{"meta":{"query_hash":"e471de49da32","filters":{"topic":"Topic Modeling"},"cohort_total":2769,"direct_labels_cover":6,"predictions_cover":2769,"exported":2769,"export_cap":100000,"truncated":false,"label_status":"direct model label, unvalidated","prediction_status":"machine_predicted_unvalidated (Codex and Gemma teacher distillation)","score_status":"score_only:v0-immature-baseline","snapshot":{"source":"OpenAlex, pinned release, all 482 partitions","release":"2026-06-24","frame_built":"2026-07-12"},"permalink":"https://metacan.xera.ac/q/e471de49da32","api":"https://metacan.xera.ac/api/v1/cohort?topic=Topic+Modeling"},"results":[{"id":"W101928771","doi":"","title":"Refining the Notions of Depth and Density in WordNet-based Semantic Similarity Measures","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"WordNet; Semantic similarity; Computer science; Artificial intelligence; Natural language processing; Similarity (geometry); Intuition; Psychology","score_opus":0.09751094935188148,"score_gpt":0.24840329148164772,"score_spread":0.15089234212976624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W101928771","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04075421,0.0016789803,0.94723415,0.0010596273,0.00018676271,0.00027420995,0.00028624263,0.00026557042,0.008260244],"genre_scores_gemma":[0.48577505,0.0012115114,0.5093975,0.00038483087,0.00022986149,0.00065209175,0.00039856762,0.00016410368,0.0017864648],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9882996,0.0046785777,0.0014497313,0.001743544,0.0033491866,0.00047928627],"domain_scores_gemma":[0.96541727,0.019544324,0.0034261413,0.005206905,0.0052192067,0.0011861821],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01076439,0.0017625819,0.0015946439,0.008668234,0.0016711965,0.005471267,0.0025862418,0.0019943395,0.0028474268],"category_scores_gemma":[0.0641271,0.001013882,0.0015836615,0.006539138,0.0060672364,0.02809148,0.009085926,0.004337406,0.00096672017],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016672531,0.00013775968,0.012004079,0.00073310954,0.0001742763,0.00018140265,0.004164037,0.016893528,0.004795886,0.77480453,0.0020986549,0.18384612],"study_design_scores_gemma":[0.00005258863,0.00026551646,0.007452593,0.0004735783,0.0001520329,0.0006149193,0.0021999055,0.11096588,0.0042100004,0.84895504,0.024462396,0.00019555564],"about_ca_topic_score_codex":0.0029649553,"about_ca_topic_score_gemma":0.0029233051,"teacher_disagreement_score":0.01076439,"about_ca_system_score_codex":0.0022790865,"about_ca_system_score_gemma":0.0017671274,"threshold_uncertainty_score":0.056928217},"labels":[],"label_agreement":null},{"id":"W104649228","doi":"","title":"Interfacing Issues for Information Extraction","year":2000,"lang":"en","type":"article","venue":"National Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Query expansion; Domain (mathematical analysis); Information extraction; Task (project management); Query language; User interface; World Wide Web; Programming language","score_opus":0.1555867381301594,"score_gpt":0.38288216419909377,"score_spread":0.22729542606893435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W104649228","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0055597844,0.0034030958,0.8729158,0.018238956,0.00086107926,0.00033345204,0.0006804127,0.009217246,0.08879017],"genre_scores_gemma":[0.21149743,0.0049548713,0.7148767,0.007915,0.001695272,0.000754842,0.004536271,0.0065025855,0.047266994],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97592133,0.012708587,0.002710472,0.0027476843,0.004931791,0.000980059],"domain_scores_gemma":[0.9300454,0.036648393,0.0017886951,0.022881823,0.0072749197,0.0013607648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02761209,0.0013585952,0.0014047226,0.004364522,0.0031895768,0.021979865,0.0058605247,0.0057422156,0.053920943],"category_scores_gemma":[0.12977554,0.0016845625,0.0022878402,0.008478053,0.0037136816,0.044439744,0.019609803,0.0054596835,0.024970872],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002866839,0.00010509473,0.0022759566,0.0010534097,0.00010813834,0.0012426123,0.0077348575,0.0021183426,0.0032621536,0.5456072,0.049837824,0.38636762],"study_design_scores_gemma":[0.000039020804,0.000046561083,0.00046355964,0.0005632852,0.00005450851,0.0010812316,0.001669349,0.012419871,0.003366397,0.5912997,0.3889163,0.00008013203],"about_ca_topic_score_codex":0.0018834046,"about_ca_topic_score_gemma":0.0010480173,"teacher_disagreement_score":0.053920943,"about_ca_system_score_codex":0.002193162,"about_ca_system_score_gemma":0.002173367,"threshold_uncertainty_score":0.1803835},"labels":[],"label_agreement":null},{"id":"W108365020","doi":"10.1007/978-3-642-25631-8_35","title":"An Aspect-Driven Random Walk Model for Topic-Focused Multi-document Summarization","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Random walk; Multi-document summarization; Exploit; Task (project management); Information retrieval; Markov chain; Artificial intelligence; Natural language processing; Machine learning","score_opus":0.03908626929132147,"score_gpt":0.2656497139802484,"score_spread":0.22656344468892692,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W108365020","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014261298,0.00095479115,0.98152816,0.00027286075,0.00011612595,0.00014197998,0.0005340617,0.0014248125,0.0007658621],"genre_scores_gemma":[0.4744334,0.0021358612,0.5023257,0.00041611725,0.00072749663,0.0008958783,0.0056863194,0.00079402403,0.012585184],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987651,0.000478048,0.00010285388,0.0003423051,0.00021858473,0.00009312177],"domain_scores_gemma":[0.9966372,0.0022771037,0.00023917247,0.00024093616,0.0005097049,0.00009587298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002680193,0.0010876423,0.0023745096,0.0016256457,0.0005464132,0.0014968216,0.0028475435,0.002161351,0.002780075],"category_scores_gemma":[0.008162238,0.0009102009,0.00172312,0.002695079,0.00047744953,0.0026652706,0.0010265274,0.0018478939,0.0018324964],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067212584,0.00023940086,0.0024267407,0.0005462951,0.0003856243,0.0002768709,0.00046275818,0.6298631,0.008950181,0.03064087,0.013303388,0.3122327],"study_design_scores_gemma":[0.000014218066,0.00003572213,0.00016076237,0.000006378153,0.000030735984,0.000025802985,0.0000062925233,0.99371064,0.00026762715,0.0050279857,0.0007029491,0.000010969574],"about_ca_topic_score_codex":0.0065625086,"about_ca_topic_score_gemma":0.009906184,"teacher_disagreement_score":0.0065625086,"about_ca_system_score_codex":0.00076464494,"about_ca_system_score_gemma":0.00090385275,"threshold_uncertainty_score":0.014174402},"labels":[],"label_agreement":null},{"id":"W10957333","doi":"","title":"Data-Driven Response Generation in Social Media","year":2011,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":580,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Phrase; Computer science; Machine translation; Natural language processing; Artificial intelligence; Social media; Speech recognition","score_opus":0.17167629064284334,"score_gpt":0.2990671782790738,"score_spread":0.12739088763623047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W10957333","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06410319,0.00017896175,0.9171577,0.00083989475,0.00021890331,0.00069386314,0.0019817452,0.011911679,0.0029140315],"genre_scores_gemma":[0.51385766,0.0001190871,0.47718588,0.00034843065,0.00016589914,0.0011839718,0.0036071553,0.0006597027,0.0028721981],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99566156,0.0030382527,0.00014228499,0.0004907762,0.0005154285,0.00015162345],"domain_scores_gemma":[0.98523456,0.011163024,0.0005277305,0.001161073,0.0016742651,0.00023947326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034953041,0.0009677841,0.0007475673,0.0008967205,0.00045443547,0.0009840966,0.0014568003,0.0011425414,0.004125992],"category_scores_gemma":[0.015798895,0.00037126135,0.00075206533,0.00093776843,0.0005710644,0.0011864309,0.0013386518,0.0012058916,0.0034532496],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020763725,0.0010793076,0.0070291823,0.0012536293,0.00025945436,0.0011831322,0.0020885048,0.16098306,0.10913193,0.015445204,0.026163863,0.6733064],"study_design_scores_gemma":[0.000096228265,0.00026260683,0.0008574524,0.000021592643,0.000026144846,0.00017769342,0.00032320776,0.9413615,0.03727432,0.012103012,0.0074446257,0.0000515788],"about_ca_topic_score_codex":0.0011620374,"about_ca_topic_score_gemma":0.0015220004,"teacher_disagreement_score":0.004125992,"about_ca_system_score_codex":0.0005362378,"about_ca_system_score_gemma":0.0007050914,"threshold_uncertainty_score":0.018485188},"labels":[],"label_agreement":null},{"id":"W116210019","doi":"","title":"Products of Hidden Markov Models.","year":2001,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Hidden Markov model; Computer science; Inference; Task (project management); Artificial intelligence; Machine learning; Markov model; Markov process; Character (mathematics); Maximum-entropy Markov model; Theoretical computer science; Markov chain; State (computer science); Range (aeronautics); Variable-order Markov model; Algorithm; Mathematics","score_opus":0.03646892649386003,"score_gpt":0.2347605916783509,"score_spread":0.19829166518449087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W116210019","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037017108,0.003117239,0.98416674,0.00066940684,0.00024194995,0.000066896166,0.00067437656,0.0011858772,0.0061758724],"genre_scores_gemma":[0.22368069,0.0062604807,0.7408358,0.0007758717,0.0006897663,0.00041581332,0.0035848618,0.0006374333,0.023119278],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99786395,0.0009482263,0.00014728046,0.00048721867,0.00043455028,0.00011877282],"domain_scores_gemma":[0.99281234,0.004840996,0.0006043708,0.0010210513,0.00053525646,0.00018590249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035443404,0.0014303208,0.0010644727,0.0014774351,0.0007298601,0.0023938767,0.0020563868,0.0019394072,0.010924007],"category_scores_gemma":[0.0140093155,0.0009893662,0.002231219,0.0018320481,0.0011123748,0.005334065,0.0021277133,0.0023879793,0.0046682125],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014498868,0.0000931985,0.0039563193,0.0005987721,0.00038034687,0.000594989,0.00062681455,0.13186471,0.0017052371,0.6492867,0.016937483,0.1938105],"study_design_scores_gemma":[0.000015149732,0.00005081729,0.00042317907,0.00007984304,0.000087321154,0.0003345983,0.000050197235,0.48146886,0.001121559,0.48593763,0.030393071,0.000037811915],"about_ca_topic_score_codex":0.0025403837,"about_ca_topic_score_gemma":0.003140166,"teacher_disagreement_score":0.010924007,"about_ca_system_score_codex":0.00088003045,"about_ca_system_score_gemma":0.0010648597,"threshold_uncertainty_score":0.036544442},"labels":[],"label_agreement":null},{"id":"W124170475","doi":"","title":"Discrepancy Between Automatic and Manual Evaluation of Summaries","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Automatic summarization; Computer science; Metric (unit); ROUGE; Natural language processing; Information retrieval; Artificial intelligence; Coherence (philosophical gambling strategy); Evaluation methods; Significant difference; Data mining; Statistics; Mathematics; Reliability engineering","score_opus":0.06725810240843816,"score_gpt":0.3314334499184848,"score_spread":0.2641753475100466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W124170475","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76173866,0.010949966,0.15686078,0.0014214134,0.0014241033,0.00091468927,0.00867538,0.032907158,0.025107784],"genre_scores_gemma":[0.9150818,0.00068000983,0.06719039,0.00045685977,0.00030641299,0.00040469007,0.00980306,0.001797898,0.0042789453],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9308292,0.04069814,0.0066364813,0.0062277056,0.014621488,0.0009869706],"domain_scores_gemma":[0.78976893,0.12448636,0.0100574745,0.026849117,0.047195874,0.0016422924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029838245,0.0015033503,0.0011952133,0.005340758,0.0009823336,0.003929817,0.0014042184,0.0014439126,0.0019234812],"category_scores_gemma":[0.14774442,0.0005157541,0.00071269553,0.00238775,0.0007076157,0.0025682037,0.0017502643,0.000891365,0.002382924],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002142732,0.00052999577,0.037098333,0.002954264,0.00095274084,0.0003980034,0.004618865,0.010001511,0.048312776,0.0026373528,0.047520828,0.84283257],"study_design_scores_gemma":[0.0009114985,0.00500102,0.40588316,0.0013397451,0.0014198192,0.0027833858,0.007247631,0.1891324,0.21510766,0.012598951,0.15734398,0.0012307783],"about_ca_topic_score_codex":0.0029128015,"about_ca_topic_score_gemma":0.0043504275,"teacher_disagreement_score":0.029838245,"about_ca_system_score_codex":0.0011516716,"about_ca_system_score_gemma":0.0007733212,"threshold_uncertainty_score":0.15780163},"labels":[],"label_agreement":null},{"id":"W127158828","doi":"","title":"Automatic Dream Sentiment Analysis","year":2006,"lang":"it","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Dream; Lexicon; Artificial intelligence; Normative; Sentiment analysis; Natural language processing; Computer science; Word (group theory); Class (philosophy); Baseline (sea); Psychology; Linguistics; Epistemology; Philosophy","score_opus":0.012854497029111579,"score_gpt":0.23832382789942536,"score_spread":0.2254693308703138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W127158828","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21919437,0.001754568,0.70625937,0.0012014005,0.00055531354,0.0014835683,0.021463068,0.013515947,0.03457233],"genre_scores_gemma":[0.5888139,0.0008664983,0.35695696,0.00022176547,0.0003537488,0.0009912595,0.036635574,0.0006451292,0.014515195],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99874216,0.00042983654,0.00014739996,0.00023045526,0.00035651663,0.00009372975],"domain_scores_gemma":[0.99737716,0.0008545646,0.00025009786,0.00028652023,0.0011476077,0.00008412864],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018670976,0.0007336641,0.00059278397,0.004055501,0.00069845596,0.0014455375,0.0007905425,0.00039331094,0.00650683],"category_scores_gemma":[0.005637248,0.00025754116,0.0007955195,0.0016159805,0.00022066357,0.0013334005,0.001013012,0.0007052913,0.0042878315],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006752183,0.00026730998,0.047266696,0.0011657568,0.0001848645,0.00027844842,0.001271164,0.00506421,0.04205946,0.010455219,0.054724827,0.8365868],"study_design_scores_gemma":[0.0001468835,0.00042493432,0.20552488,0.00039328533,0.00030624188,0.0015810337,0.0049987277,0.48395813,0.08294401,0.041621394,0.17777808,0.00032231194],"about_ca_topic_score_codex":0.0021475446,"about_ca_topic_score_gemma":0.0032982358,"teacher_disagreement_score":0.00650683,"about_ca_system_score_codex":0.00063374935,"about_ca_system_score_gemma":0.0006116897,"threshold_uncertainty_score":0.021767497},"labels":[],"label_agreement":null},{"id":"W133424112","doi":"10.1007/978-3-642-30353-1_29","title":"Text Similarity Using Google Tri-grams","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Similarity (geometry); Artificial intelligence; Set (abstract data type); Natural language processing; Word (group theory); n-gram; Data set; Information retrieval; Language model; Mathematics; Image (mathematics)","score_opus":0.049105195206289545,"score_gpt":0.27087556995842194,"score_spread":0.2217703747521324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W133424112","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.171812,0.011148781,0.66584,0.0009222667,0.0027613724,0.0015801417,0.036398064,0.06532615,0.044211324],"genre_scores_gemma":[0.4678721,0.0024752447,0.4484072,0.00018109268,0.0011886642,0.00076589826,0.049754936,0.0031069985,0.026247947],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99774927,0.0002799762,0.00030029126,0.00048127913,0.0010252801,0.00016388312],"domain_scores_gemma":[0.9974195,0.0006572529,0.0002639217,0.00039987257,0.00109739,0.00016203367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009729846,0.0012431382,0.0016873322,0.017130395,0.0011987443,0.0028294486,0.0010854101,0.0013202808,0.01576247],"category_scores_gemma":[0.0061780056,0.00038910614,0.0016765633,0.012994833,0.00034456662,0.004489659,0.0021380933,0.0008303757,0.015104949],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080610695,0.00022332647,0.0049121473,0.0009114697,0.00029006763,0.00025742548,0.00036760856,0.0039869524,0.020299828,0.006310122,0.05811567,0.90351933],"study_design_scores_gemma":[0.00032454092,0.0011565057,0.039066244,0.00041537176,0.0007221529,0.0049062436,0.0025173384,0.6598907,0.06802795,0.07146294,0.15112695,0.00038297087],"about_ca_topic_score_codex":0.0032340672,"about_ca_topic_score_gemma":0.005365049,"teacher_disagreement_score":0.017130395,"about_ca_system_score_codex":0.0005947514,"about_ca_system_score_gemma":0.0009227031,"threshold_uncertainty_score":0.05273074},"labels":[],"label_agreement":null},{"id":"W134362049","doi":"","title":"Modeling Language Acquisition at Multiple Temporal Scales","year":2000,"lang":"en","type":"article","venue":"eScholarship (California Digital Library)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McMaster University","funders":"","keywords":"Verb; Context (archaeology); Representation (politics); Linguistics; Computer science; Noun; Argument (complex analysis); Artificial intelligence; Natural language processing; Cognitive science; Psychology; Cognitive psychology; History","score_opus":0.013881826250613445,"score_gpt":0.21217777493749737,"score_spread":0.19829594868688394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W134362049","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40317315,0.0013298054,0.5671724,0.0024613063,0.000106725805,0.00007492931,0.0006080853,0.00073305145,0.024340527],"genre_scores_gemma":[0.9198239,0.0006461257,0.06749877,0.0001236274,0.000053064745,0.00012878014,0.0002581122,0.00013045134,0.011337215],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99963737,0.000121676305,0.000019230665,0.00009609154,0.00005838349,0.00006732889],"domain_scores_gemma":[0.9979887,0.0012377707,0.00027107683,0.00012690232,0.00019955725,0.00017600271],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001052745,0.00048515556,0.0004951297,0.00068065524,0.0003783009,0.0017113637,0.0010315208,0.001183826,0.0037128048],"category_scores_gemma":[0.005726404,0.00069409807,0.0010334841,0.0006305141,0.0009887966,0.002902291,0.0017667338,0.0016839787,0.00054108247],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011257204,0.000069276015,0.006819068,0.00006145938,0.000085192514,0.00021960263,0.00033849498,0.87597334,0.0025168296,0.0916872,0.0009475517,0.02116945],"study_design_scores_gemma":[0.000007288006,0.000013227959,0.0007065665,0.000006627634,0.000010529426,0.00002153744,0.000027374488,0.97593963,0.00013865263,0.022416802,0.0007028289,0.000008809712],"about_ca_topic_score_codex":0.024522465,"about_ca_topic_score_gemma":0.019104995,"teacher_disagreement_score":0.024522465,"about_ca_system_score_codex":0.0018593904,"about_ca_system_score_gemma":0.00117066,"threshold_uncertainty_score":0.04875946},"labels":[],"label_agreement":null},{"id":"W13618705","doi":"10.1007/978-3-642-38457-8_13","title":"Feature Combination for Sentence Similarity","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; SemEval; Sentence; Artificial intelligence; Task (project management); Similarity (geometry); Feature (linguistics); Feature vector; Ranking (information retrieval); Semantic similarity; Natural language processing; Pattern recognition (psychology); Support vector machine; Semantic feature","score_opus":0.024271195096126942,"score_gpt":0.25115860976537,"score_spread":0.2268874146692431,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W13618705","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07616665,0.0038990257,0.89749765,0.00032785573,0.0007447952,0.00029714557,0.0035603968,0.009641966,0.007864437],"genre_scores_gemma":[0.5764087,0.0012578934,0.39369172,0.00021662508,0.0009008682,0.00057052594,0.0149179725,0.0008942524,0.011141348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983981,0.00034758358,0.00014970914,0.00044528287,0.0004930993,0.00016613968],"domain_scores_gemma":[0.99810225,0.00069692545,0.00010426476,0.0003451578,0.0006446078,0.00010684243],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001379226,0.0013112874,0.0019254284,0.003249906,0.0006620964,0.0012075561,0.0011870435,0.0011042925,0.008699953],"category_scores_gemma":[0.0038025933,0.0003400565,0.0015232618,0.0030996378,0.00026942085,0.0020783986,0.0016246607,0.0011515938,0.005576087],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089373195,0.00023333963,0.0014696483,0.00023807873,0.00019021412,0.00012974508,0.0000592908,0.0059422194,0.033692636,0.0016851387,0.017729474,0.9377364],"study_design_scores_gemma":[0.00014990561,0.0009497169,0.012101785,0.00008960848,0.0006295897,0.00095050014,0.00016888823,0.8911175,0.05822602,0.01593244,0.019532768,0.0001511671],"about_ca_topic_score_codex":0.001404504,"about_ca_topic_score_gemma":0.0017282129,"teacher_disagreement_score":0.008699953,"about_ca_system_score_codex":0.0003969536,"about_ca_system_score_gemma":0.000595253,"threshold_uncertainty_score":0.029104233},"labels":[],"label_agreement":null},{"id":"W136188959","doi":"","title":"CLASSY and TAC 2008 Metrics.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval","score_opus":0.016530037573413182,"score_gpt":0.2413698430885062,"score_spread":0.22483980551509303,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W136188959","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45473355,0.011709923,0.06727741,0.007987422,0.0076804813,0.007071273,0.22939959,0.062003404,0.15213703],"genre_scores_gemma":[0.5421538,0.0007343928,0.100945316,0.0009810027,0.0014337722,0.004486065,0.30239195,0.005888921,0.040984925],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9645046,0.012067935,0.0028441136,0.0030689957,0.015379717,0.002134753],"domain_scores_gemma":[0.9362193,0.018076593,0.005317517,0.0091625005,0.026681876,0.004542323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014776312,0.002688599,0.0017319856,0.010867146,0.0022937593,0.0034028292,0.0019914461,0.0026560277,0.009741705],"category_scores_gemma":[0.07246227,0.00043302405,0.0011013624,0.0067129,0.00064347923,0.0037911837,0.002674523,0.0020416772,0.005759602],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015299569,0.0015789416,0.019393828,0.0015036257,0.00044365294,0.00013715889,0.0006570267,0.0054636616,0.0060105696,0.0029513636,0.77677333,0.18355684],"study_design_scores_gemma":[0.0020883037,0.0046894923,0.19483837,0.0004490366,0.0006526631,0.0011575748,0.002060105,0.20193797,0.035504445,0.0090053165,0.54695386,0.0006628636],"about_ca_topic_score_codex":0.020014012,"about_ca_topic_score_gemma":0.038059194,"teacher_disagreement_score":0.020014012,"about_ca_system_score_codex":0.002962271,"about_ca_system_score_gemma":0.0023853444,"threshold_uncertainty_score":0.078145504},"labels":[],"label_agreement":null},{"id":"W137601953","doi":"","title":"Unsupervised Part-of-Speech Tagging in Noisy and Esoteric Domains With a Syntactic-Semantic Bayesian HMM","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Hidden Markov model; Computer science; Artificial intelligence; Natural language processing; Bayesian probability; Generative grammar; Unsupervised learning; Topic model; Probabilistic logic; Generative model; Speech recognition","score_opus":0.015672723018685213,"score_gpt":0.2306368129663256,"score_spread":0.21496408994764038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W137601953","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0352426,0.0001671136,0.9614942,0.00031069628,0.000048271948,0.00004510525,0.00047421653,0.0009096391,0.0013082629],"genre_scores_gemma":[0.6103826,0.00040416236,0.37943247,0.00025781474,0.00015006126,0.00030295228,0.002926107,0.00043059225,0.005713192],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99872845,0.0006528838,0.00006873251,0.0003102563,0.0001396184,0.00010007056],"domain_scores_gemma":[0.994406,0.0043134317,0.000249985,0.00051293516,0.0003880164,0.00012960668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025914072,0.0006164481,0.00094678055,0.0011947368,0.00095003954,0.0014563625,0.0015029309,0.0015054088,0.00166305],"category_scores_gemma":[0.006293384,0.00080924097,0.0010320048,0.0014332877,0.0011282052,0.0023528389,0.0018320209,0.0019564768,0.0021920174],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010005371,0.00044724045,0.020091442,0.0003821527,0.00027593996,0.0012427198,0.0027220263,0.533384,0.04107698,0.07905506,0.01332239,0.30699956],"study_design_scores_gemma":[0.00001009599,0.0000152020975,0.0012746933,0.000013253663,0.000015506206,0.00007229726,0.000055493485,0.97465384,0.0021773404,0.020505827,0.0011790808,0.000027266506],"about_ca_topic_score_codex":0.0066787843,"about_ca_topic_score_gemma":0.013601041,"teacher_disagreement_score":0.0066787843,"about_ca_system_score_codex":0.00094569975,"about_ca_system_score_gemma":0.0011424534,"threshold_uncertainty_score":0.013704777},"labels":[],"label_agreement":null},{"id":"W138032095","doi":"10.1609/aiide.v5i1.12371","title":"Supporting Dialogue Generation for Story-Based Games","year":2009,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence and Interactive Digital Entertainment","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Immersion (mathematics); Character (mathematics); Game design; Human–computer interaction; Multimedia; Artificial intelligence; Mathematics","score_opus":0.063535911672882,"score_gpt":0.31215809089826363,"score_spread":0.24862217922538163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W138032095","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03795336,0.0002508951,0.94931376,0.00020449459,0.00006717494,0.00021762565,0.00012414537,0.0060958182,0.0057726954],"genre_scores_gemma":[0.6456508,0.00021998417,0.34932968,0.00013035377,0.000063217616,0.00040909986,0.00062396406,0.0005787788,0.0029941364],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99845135,0.0008723645,0.00009877076,0.00020282243,0.0002645341,0.00011008938],"domain_scores_gemma":[0.99654096,0.002520504,0.00016608513,0.00029833836,0.0002452169,0.00022888141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015986725,0.0008950589,0.00043823896,0.0005487799,0.0006441535,0.0018627617,0.0014466267,0.001085284,0.0049643666],"category_scores_gemma":[0.008970414,0.00046465354,0.00051452336,0.00024746094,0.00069859816,0.0024538531,0.0026217788,0.0010682147,0.0016477478],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018952956,0.0013653378,0.0061305,0.0016699363,0.00022608192,0.0020718477,0.015809486,0.112000465,0.14057185,0.120523915,0.0202535,0.57748187],"study_design_scores_gemma":[0.00015007248,0.00024597606,0.0009400683,0.000116732655,0.00005373357,0.00052133953,0.00094574824,0.85140455,0.032325692,0.05272609,0.0605017,0.000068292255],"about_ca_topic_score_codex":0.00054606877,"about_ca_topic_score_gemma":0.00070246554,"teacher_disagreement_score":0.0049643666,"about_ca_system_score_codex":0.0003467078,"about_ca_system_score_gemma":0.000406213,"threshold_uncertainty_score":0.016607523},"labels":[],"label_agreement":null},{"id":"W1422654747","doi":"10.1007/978-3-642-54943-4_12","title":"Measuring Conceptual Entanglement in Collections of Documents","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; University of British Columbia","funders":"","keywords":"Quantum entanglement; Computer science; Context (archaeology); Phenomenon; Conceptual model; Natural language; Natural language processing; Artificial intelligence; Quantum; Quantum mechanics; Physics","score_opus":0.03580766092353518,"score_gpt":0.24960816396945096,"score_spread":0.21380050304591577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1422654747","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.585288,0.0061100665,0.39098454,0.0006493583,0.00014089876,0.00022113277,0.0025295967,0.0013011874,0.012775183],"genre_scores_gemma":[0.9185558,0.0017064679,0.07219745,0.00008629599,0.0001544434,0.00021422986,0.004167764,0.00032635545,0.0025912707],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994179,0.0019815373,0.00045367138,0.0011226017,0.0018249032,0.00043824958],"domain_scores_gemma":[0.974678,0.017204037,0.0021490366,0.0031993585,0.0017002411,0.0010693619],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004064492,0.00074508437,0.0013966085,0.00761355,0.0021876092,0.006738798,0.0015236521,0.0014428048,0.0037138718],"category_scores_gemma":[0.035442885,0.0010274603,0.0011660035,0.011154229,0.0023742397,0.012864668,0.006656917,0.0022367046,0.0008573021],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002871947,0.0010686418,0.10557759,0.0022861676,0.0015673212,0.00079577824,0.01069716,0.052704524,0.049526274,0.2753447,0.011380022,0.4861799],"study_design_scores_gemma":[0.000097852724,0.00044441962,0.059027355,0.00033843823,0.0007242249,0.0010501414,0.0057979696,0.29701802,0.02705778,0.588873,0.019293927,0.00027685714],"about_ca_topic_score_codex":0.0019758022,"about_ca_topic_score_gemma":0.0021128312,"teacher_disagreement_score":0.00761355,"about_ca_system_score_codex":0.0016227391,"about_ca_system_score_gemma":0.0009785817,"threshold_uncertainty_score":0.021495283},"labels":[],"label_agreement":null},{"id":"W142564969","doi":"10.1136/amiajnl-2013-001624","title":"À la Recherche du Temps Perdu: extracting temporal relations from medical text in the 2012 i2b2 NLP challenge","year":2013,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Topic Modeling","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"National Research Council Canada","funders":"U.S. National Library of Medicine","keywords":"Recall; Natural language processing; Artificial intelligence; Computer science; Task (project management); Relation (database); Relationship extraction; Post hoc; Information retrieval; Psychology; Information extraction; Cognitive psychology; Data mining; Medicine; Engineering","score_opus":0.08013885544504842,"score_gpt":0.3295911244913692,"score_spread":0.24945226904632076,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W142564969","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31593913,0.024936259,0.39005795,0.041949242,0.0053976653,0.0025318284,0.13888004,0.03895889,0.041349012],"genre_scores_gemma":[0.34260166,0.0039528273,0.4056005,0.002790432,0.0018818869,0.0016656582,0.21767254,0.0029724974,0.020862041],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99019927,0.004174786,0.0009488626,0.0020623815,0.002169936,0.00044482303],"domain_scores_gemma":[0.97633696,0.016175494,0.0011967281,0.0018572506,0.0033781254,0.0010553832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012444331,0.0018341264,0.0012612579,0.0043521444,0.002415804,0.003814273,0.0014875344,0.003142558,0.004367372],"category_scores_gemma":[0.041545983,0.00058590755,0.0013612983,0.003889744,0.0010291052,0.0038819162,0.0031741355,0.0027252561,0.004044179],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013522542,0.00047508735,0.018140959,0.004711219,0.00037515716,0.0017758238,0.0072024097,0.014958721,0.022363072,0.009746752,0.45181975,0.4670789],"study_design_scores_gemma":[0.00050828035,0.00076740846,0.049577836,0.0010093239,0.00037354176,0.004007031,0.0072187595,0.21358635,0.04836629,0.018428547,0.6556517,0.0005050021],"about_ca_topic_score_codex":0.026740681,"about_ca_topic_score_gemma":0.028653763,"teacher_disagreement_score":0.026740681,"about_ca_system_score_codex":0.00292027,"about_ca_system_score_gemma":0.00652619,"threshold_uncertainty_score":0.06581271},"labels":[],"label_agreement":null},{"id":"W1436680051","doi":"10.3115/v1/w15-0609","title":"Towards Automatic Description of Knowledge Components","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science","score_opus":0.12789403296203156,"score_gpt":0.28945505554558926,"score_spread":0.1615610225835577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1436680051","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021745354,0.0017705793,0.94624,0.00087045017,0.000096696065,0.0003487829,0.015908709,0.01015431,0.0028651436],"genre_scores_gemma":[0.22345667,0.001289944,0.7053884,0.00038861085,0.00022586831,0.00092562725,0.06404146,0.0010152395,0.0032681262],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99578226,0.0010500592,0.00048529916,0.0012367471,0.0011452896,0.00030042868],"domain_scores_gemma":[0.98808575,0.0069965,0.00070501334,0.0019371418,0.001931161,0.0003445102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027794484,0.0021880944,0.0016558613,0.009686606,0.00094046077,0.0043225368,0.0036614258,0.0028050467,0.004102757],"category_scores_gemma":[0.023441361,0.0009905867,0.0027325999,0.008397498,0.000938479,0.0062201396,0.003599127,0.0038538147,0.0044799936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053044147,0.00036893468,0.013812712,0.0016930744,0.0003612279,0.00063616474,0.0016000486,0.035667103,0.012780532,0.036180273,0.059023447,0.83734614],"study_design_scores_gemma":[0.0000907108,0.000070520626,0.0042175627,0.0002655906,0.00020608668,0.0004870829,0.00051344524,0.8244298,0.009814801,0.122827254,0.03699238,0.00008479303],"about_ca_topic_score_codex":0.008569335,"about_ca_topic_score_gemma":0.009518312,"teacher_disagreement_score":0.009686606,"about_ca_system_score_codex":0.001917481,"about_ca_system_score_gemma":0.0036806013,"threshold_uncertainty_score":0.017038882},"labels":[],"label_agreement":null},{"id":"W14532928","doi":"10.1155/2003/213213","title":"Mining and re-ranking for answering biographical queries on the web","year":2006,"lang":"en","type":"article","venue":"National Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Association of Gastroenterology","keywords":"Computer science; Bootstrapping (finance); Ranking (information retrieval); Information retrieval; Question answering; Knowledge base; Web mining; World Wide Web; Web page; Artificial intelligence","score_opus":0.14680838160324203,"score_gpt":0.3276615497888354,"score_spread":0.18085316818559335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W14532928","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6062796,0.014956821,0.28202617,0.008261935,0.0010220971,0.0021592907,0.04962766,0.024458507,0.011207922],"genre_scores_gemma":[0.6123709,0.0030092301,0.30928373,0.00054802204,0.00075935817,0.00046776922,0.06746802,0.0005059808,0.0055870744],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963691,0.0011208985,0.0005311875,0.00054076896,0.0010511801,0.0003868186],"domain_scores_gemma":[0.99225205,0.0038006296,0.0006012669,0.0012302161,0.001686808,0.00042912952],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002314309,0.0016475024,0.0020717275,0.016014917,0.0010752662,0.0027497264,0.0018830424,0.0019658464,0.0036384112],"category_scores_gemma":[0.013191662,0.00047638488,0.0015316235,0.008544019,0.00040261206,0.0043990207,0.0010665349,0.0014125307,0.003935015],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084176986,0.0017795371,0.058671966,0.0017325116,0.00058246637,0.0009834166,0.0008094574,0.016329627,0.030163554,0.0035266741,0.07523067,0.8093482],"study_design_scores_gemma":[0.00016774052,0.0009074001,0.05641776,0.00029905635,0.00060397567,0.0022508898,0.0025057453,0.8335982,0.029899348,0.02363462,0.049512304,0.00020284511],"about_ca_topic_score_codex":0.012740533,"about_ca_topic_score_gemma":0.025428921,"teacher_disagreement_score":0.016014917,"about_ca_system_score_codex":0.00084765133,"about_ca_system_score_gemma":0.0014469747,"threshold_uncertainty_score":0.025332749},"labels":[],"label_agreement":null},{"id":"W1481096275","doi":"10.34105/j.kmel.2014.06.002","title":"Use of global context for handling noisy names in discussion texts of a homeopathy discussion forum","year":2014,"lang":"en","type":"article","venue":"Knowledge Management & E-Learning An International Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University","keywords":"Context (archaeology); Computer science; Task (project management); Homeopathy; Natural language processing; Artificial intelligence; Information retrieval; Linguistics; History; Medicine; Engineering","score_opus":0.02262400135493118,"score_gpt":0.29429362097023715,"score_spread":0.27166961961530595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1481096275","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40907168,0.0042152912,0.5678216,0.0010635344,0.00050688063,0.00062780315,0.0010373809,0.004492682,0.011163152],"genre_scores_gemma":[0.8648143,0.00050561124,0.13239272,0.00010035722,0.00019897848,0.00014408545,0.00054903975,0.00013858783,0.0011563087],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964896,0.0017374187,0.00026877876,0.0007654925,0.0005605534,0.00017811486],"domain_scores_gemma":[0.9912299,0.004718135,0.0011521556,0.00081181823,0.0016827963,0.00040530704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034428746,0.0008296517,0.00091382227,0.0039531514,0.0018634392,0.002741054,0.0007595859,0.000995011,0.0014733283],"category_scores_gemma":[0.013139419,0.00033891335,0.00046584563,0.0020003126,0.0008126055,0.0058944654,0.002133589,0.0009309172,0.0005633316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015405064,0.00035304067,0.051141143,0.0014822877,0.000248584,0.002110217,0.018926078,0.009936255,0.0749119,0.014173363,0.0050542266,0.8201224],"study_design_scores_gemma":[0.00025118756,0.0016457195,0.1520906,0.0009162167,0.001671678,0.004686942,0.021382023,0.5014913,0.13220362,0.05454838,0.12804303,0.0010692757],"about_ca_topic_score_codex":0.0024698349,"about_ca_topic_score_gemma":0.0045394464,"teacher_disagreement_score":0.0039531514,"about_ca_system_score_codex":0.00055949384,"about_ca_system_score_gemma":0.00107983,"threshold_uncertainty_score":0.018207848},"labels":[],"label_agreement":null},{"id":"W1492071093","doi":"10.48550/arxiv.1401.0509","title":"Zero-Shot Learning for Semantic Utterance Classification","year":2013,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Semantic space; Utterance; Discriminative model; Zero (linguistics); Classifier (UML); Natural language processing; Training set; Machine learning; Linguistics","score_opus":0.14724743541191562,"score_gpt":0.21097881591590373,"score_spread":0.0637313805039881,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1492071093","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032816056,0.0011317998,0.9598891,0.00043484409,0.0001289327,0.00013315673,0.0006660214,0.0025908316,0.002209416],"genre_scores_gemma":[0.72837436,0.0004917555,0.25990462,0.00043379026,0.00031473403,0.0003641161,0.0047460264,0.00031550953,0.0050550974],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99796355,0.00068151864,0.00011232395,0.0006062969,0.0004313079,0.00020503729],"domain_scores_gemma":[0.99780554,0.001200489,0.00012796395,0.00041727914,0.00033045342,0.00011826144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017660513,0.0014425294,0.0019537478,0.0018016461,0.00086869916,0.0011342425,0.003743962,0.0017962577,0.0032148473],"category_scores_gemma":[0.0059118224,0.00050068914,0.001114744,0.0016245091,0.0013648559,0.0036355653,0.002140896,0.002846696,0.0012385094],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008367273,0.0008546379,0.0053883274,0.0006659647,0.0002823953,0.00029745052,0.0007548219,0.09828565,0.013352948,0.037073087,0.024632545,0.8175754],"study_design_scores_gemma":[0.000023274208,0.000081313876,0.0005802715,0.000022238692,0.000025978623,0.00007074232,0.0001024157,0.9588824,0.004390368,0.03371713,0.002079418,0.000024538685],"about_ca_topic_score_codex":0.007619627,"about_ca_topic_score_gemma":0.007566341,"teacher_disagreement_score":0.007619627,"about_ca_system_score_codex":0.0013893805,"about_ca_system_score_gemma":0.001408091,"threshold_uncertainty_score":0.015150547},"labels":[],"label_agreement":null},{"id":"W1498740961","doi":"10.48550/arxiv.1410.0718","title":"Not All Neural Embeddings are Born Equal","year":2014,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Natural language processing; Machine translation; Artificial intelligence; Similarity (geometry); Word (group theory); Translation (biology); Artificial neural network; Linguistics; Image (mathematics)","score_opus":0.10884516484628487,"score_gpt":0.20611422685168357,"score_spread":0.0972690620053987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1498740961","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14911264,0.012935239,0.6843959,0.056884054,0.0032041732,0.00005694206,0.0026581637,0.001533936,0.08921894],"genre_scores_gemma":[0.8870658,0.0076180943,0.068243705,0.0055549997,0.0009757438,0.00011474897,0.002237683,0.0007715311,0.027417582],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998694,0.0004427138,0.000063115986,0.00047146238,0.0002250707,0.00010375754],"domain_scores_gemma":[0.9969132,0.0013751674,0.0002022089,0.0010038023,0.00035089228,0.00015457168],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014771713,0.00069641316,0.0007627015,0.0007717178,0.00078292115,0.004251136,0.0009764475,0.001820196,0.0068508144],"category_scores_gemma":[0.016521519,0.00061769463,0.00047465492,0.0010355894,0.002939134,0.019013632,0.0024162817,0.0037689607,0.0036420915],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016378253,0.000048949783,0.00379801,0.0003333582,0.00036711412,0.00015536431,0.0004535983,0.018872378,0.0027980262,0.70769656,0.028049994,0.23726287],"study_design_scores_gemma":[0.000012884626,0.0000215571,0.0009049763,0.000064131345,0.00004578892,0.0001502683,0.00017137511,0.022824645,0.0011292049,0.9561597,0.01849026,0.000025208114],"about_ca_topic_score_codex":0.001031255,"about_ca_topic_score_gemma":0.0010846484,"teacher_disagreement_score":0.0068508144,"about_ca_system_score_codex":0.0007312734,"about_ca_system_score_gemma":0.00044331164,"threshold_uncertainty_score":0.022918224},"labels":[],"label_agreement":null},{"id":"W1503829033","doi":"","title":"Experiments with the Negotiated Boolean Queries of the TREC 2008 Legal Track","year":2008,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Relevance (law); Task (project management); Computer science; Relevance feedback; Information retrieval; Recall; Standard Boolean model; Rank (graph theory); Boolean expression; Learning to rank; Boolean network; And-inverter graph; Theoretical computer science; Boolean function; Data mining; Artificial intelligence; Algorithm; Mathematics; Ranking (information retrieval); Combinatorics; Psychology; Cognitive psychology","score_opus":0.04573418768938414,"score_gpt":0.24959330938408245,"score_spread":0.2038591216946983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1503829033","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97261727,0.0022498874,0.009712175,0.00072923914,0.00038888768,0.0008673828,0.0031922988,0.004077038,0.0061658733],"genre_scores_gemma":[0.93731666,0.00046641123,0.040339682,0.0004498585,0.00041174487,0.0007778126,0.014030347,0.0006252815,0.0055822595],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9843163,0.008304994,0.0013671031,0.002150874,0.0032294388,0.00063130556],"domain_scores_gemma":[0.9626464,0.028518025,0.001482786,0.002916577,0.0032748391,0.0011613291],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0119051,0.0019263175,0.0022097435,0.0013375221,0.0020686362,0.0019956194,0.0022434997,0.0024166508,0.0038093368],"category_scores_gemma":[0.040022634,0.00073528895,0.0009949096,0.0024220021,0.0010840874,0.0028515118,0.0017639139,0.002137366,0.0016312313],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.032369025,0.02388874,0.038365487,0.008505081,0.00380754,0.0024054607,0.0058686202,0.1216327,0.16252884,0.0032088214,0.10051386,0.4969058],"study_design_scores_gemma":[0.008120317,0.03329992,0.117036976,0.0003721862,0.0019065274,0.0031567623,0.0037612566,0.5885445,0.18037146,0.004405428,0.058043182,0.0009815376],"about_ca_topic_score_codex":0.014327412,"about_ca_topic_score_gemma":0.014742744,"teacher_disagreement_score":0.014327412,"about_ca_system_score_codex":0.001887333,"about_ca_system_score_gemma":0.0015692704,"threshold_uncertainty_score":0.06296098},"labels":[],"label_agreement":null},{"id":"W150418520","doi":"10.1007/3-540-47922-8_30","title":"Retrieval of Short Documents from Discussion Forums","year":2002,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Recall; Information retrieval; Precision and recall; Document retrieval; World Wide Web","score_opus":0.021740752960943487,"score_gpt":0.25364779628556444,"score_spread":0.23190704332462095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W150418520","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4529953,0.040584493,0.3905641,0.004214181,0.004087874,0.0021587608,0.042928215,0.016325427,0.04614151],"genre_scores_gemma":[0.506551,0.01602828,0.2863444,0.0005300284,0.0045362134,0.001403887,0.09653787,0.0016861442,0.086382195],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989121,0.0003049551,0.00009979558,0.00014951361,0.00042580525,0.00010780078],"domain_scores_gemma":[0.9951084,0.0030711433,0.0002759151,0.0002941934,0.0009526559,0.0002976275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014934259,0.0015782635,0.0017521849,0.011690983,0.0016157636,0.0030083077,0.0010494532,0.0013002186,0.014673841],"category_scores_gemma":[0.0075184866,0.0004928755,0.0009458835,0.007967322,0.0003097098,0.004065681,0.0015977923,0.00068682793,0.010279949],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015918257,0.0005366738,0.00315571,0.0027629475,0.00019072974,0.00069304364,0.0016305605,0.002771655,0.054836534,0.0043126615,0.1014162,0.8261014],"study_design_scores_gemma":[0.0012939858,0.0034150751,0.040049404,0.002135714,0.0017357355,0.005655777,0.0073055243,0.20005107,0.18386674,0.056224536,0.4976655,0.00060106296],"about_ca_topic_score_codex":0.0009600931,"about_ca_topic_score_gemma":0.0014787208,"teacher_disagreement_score":0.014673841,"about_ca_system_score_codex":0.0006337306,"about_ca_system_score_gemma":0.00091454183,"threshold_uncertainty_score":0.049088836},"labels":[],"label_agreement":null},{"id":"W1505951874","doi":"10.48550/arxiv.cs/0412024","title":"Human-Level Performance on Word Analogy Questions by Latent Relational Analysis","year":2004,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Analogy; Word (group theory); Computer science; Natural language processing; Linguistics; Psychology; Cognitive psychology; Philosophy","score_opus":0.10671591628618171,"score_gpt":0.2968550252287323,"score_spread":0.1901391089425506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1505951874","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8705073,0.0039679264,0.07553446,0.0008104858,0.0003869638,0.0006948,0.0032927673,0.009993534,0.034811735],"genre_scores_gemma":[0.96397024,0.00049756665,0.025339946,0.0002466528,0.00011679299,0.00033206417,0.0049408083,0.0003825691,0.004173423],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97326815,0.0155404955,0.0013167021,0.0054731136,0.003729731,0.00067166693],"domain_scores_gemma":[0.9379063,0.050547726,0.0024156563,0.004829553,0.0029998703,0.0013008313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01301709,0.001953546,0.0018754873,0.0036099898,0.0008885768,0.0039842906,0.0016420502,0.0027688134,0.009102009],"category_scores_gemma":[0.07121853,0.0004150787,0.0015122866,0.002476372,0.0010357236,0.008162241,0.0042269467,0.0015289977,0.007107525],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004123202,0.002117258,0.08665517,0.0022147065,0.0010632857,0.0005528625,0.0061712954,0.029884182,0.02689528,0.004956307,0.032128572,0.80323786],"study_design_scores_gemma":[0.00078663847,0.005242489,0.240819,0.0003880256,0.0005646084,0.001863149,0.0051404843,0.63268286,0.034410335,0.03386372,0.04345235,0.00078629644],"about_ca_topic_score_codex":0.003741716,"about_ca_topic_score_gemma":0.0022737603,"teacher_disagreement_score":0.01301709,"about_ca_system_score_codex":0.00096973666,"about_ca_system_score_gemma":0.0007400933,"threshold_uncertainty_score":0.068841755},"labels":[],"label_agreement":null},{"id":"W1507075132","doi":"10.1002/asi.23216","title":"Truth and deception at the rhetorical structure level","year":2014,"lang":"en","type":"article","venue":"Journal of the Association for Information Science and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":144,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Deception; Rhetorical question; Computer science; Robustness (evolution); Coherence (philosophical gambling strategy); Sample (material); Psychology; Social psychology; Artificial intelligence; Epistemology; Linguistics; Mathematics; Statistics","score_opus":0.01646754413629537,"score_gpt":0.25067359460632643,"score_spread":0.23420605047003107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1507075132","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4282879,0.003603006,0.5377385,0.0028481472,0.0001726498,0.00038290667,0.0010045792,0.0004705531,0.025491733],"genre_scores_gemma":[0.93092084,0.0004998554,0.06689337,0.00013190709,0.00011244285,0.00014701972,0.0004197339,0.000057679383,0.0008172329],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9864174,0.0074254633,0.0011620879,0.001205159,0.0035072828,0.0002826331],"domain_scores_gemma":[0.7596351,0.2016709,0.019624183,0.010308808,0.007889626,0.00087126513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014380714,0.0006337974,0.00077166484,0.0093440665,0.001950419,0.006271641,0.0011039942,0.0015260312,0.0045515336],"category_scores_gemma":[0.13215375,0.0004612232,0.0005924349,0.0035137716,0.0046809316,0.011337028,0.0031392705,0.0018879853,0.0005854908],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010463839,0.0002910657,0.12871534,0.0027507066,0.00035233665,0.00069206685,0.07579021,0.012213287,0.024972964,0.28580365,0.0021760315,0.465196],"study_design_scores_gemma":[0.00010358484,0.00083188043,0.18521996,0.002206509,0.000382825,0.0026874845,0.035658997,0.16638118,0.038391035,0.5210581,0.046658292,0.00042021312],"about_ca_topic_score_codex":0.0012695277,"about_ca_topic_score_gemma":0.0012589763,"teacher_disagreement_score":0.014380714,"about_ca_system_score_codex":0.0017355302,"about_ca_system_score_gemma":0.0013334963,"threshold_uncertainty_score":0.07605338},"labels":[],"label_agreement":null},{"id":"W1510000161","doi":"","title":"A Probabilistic Answer Type Model","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Type (biology); Set (abstract data type); Probabilistic logic; Artificial intelligence; Machine learning; Questions and answers; Answer set programming; Natural language processing; Programming language","score_opus":0.02411001214264613,"score_gpt":0.23279789367626263,"score_spread":0.2086878815336165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1510000161","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0074621784,0.0002537996,0.978621,0.0016547769,0.00012274853,0.00034636472,0.0021638519,0.0018120283,0.007563184],"genre_scores_gemma":[0.35471585,0.00066645414,0.61011356,0.001419932,0.000577323,0.002347743,0.006613829,0.000541149,0.0230042],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916054,0.003034591,0.00062155176,0.0019939197,0.0023171026,0.00042739083],"domain_scores_gemma":[0.98005277,0.01326154,0.00096817926,0.0019166927,0.0033282214,0.00047265211],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007812647,0.0009956018,0.0013316801,0.0042607547,0.0010669327,0.0045422507,0.004209879,0.0041952343,0.0127555365],"category_scores_gemma":[0.03724326,0.0011771444,0.0028168664,0.0032863163,0.0016933316,0.009307625,0.0021021883,0.0031594245,0.005088226],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005465083,0.00029413163,0.007543646,0.000621185,0.00023848737,0.00048861635,0.0016047416,0.13934898,0.0051972466,0.61569685,0.029931383,0.19848818],"study_design_scores_gemma":[0.00008402022,0.00007140897,0.0009043105,0.00007565083,0.0001164067,0.00042448883,0.000105696774,0.6976039,0.0013498443,0.27969673,0.0194988,0.0000687053],"about_ca_topic_score_codex":0.004542932,"about_ca_topic_score_gemma":0.004939827,"teacher_disagreement_score":0.0127555365,"about_ca_system_score_codex":0.0019367131,"about_ca_system_score_gemma":0.0024575852,"threshold_uncertainty_score":0.04267156},"labels":[],"label_agreement":null},{"id":"W1510597680","doi":"","title":"Using learned browsing behavior models to recommend relevant web pages","year":2005,"lang":"en","type":"article","venue":"Alexandria (UniSG) (University of St.Gallen)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Session (web analytics); Computer science; World Wide Web; Web page; Set (abstract data type); Information retrieval; Web navigation; User information; Empirical research; Information system; Mathematics; Engineering","score_opus":0.11982968127798144,"score_gpt":0.26351996843197817,"score_spread":0.14369028715399673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1510597680","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45279577,0.0013405083,0.5384443,0.0008289493,0.000061603954,0.00021873099,0.0014329354,0.0027353007,0.0021419148],"genre_scores_gemma":[0.8701752,0.00046166786,0.12341357,0.00017622212,0.0001024824,0.00022480912,0.0036651883,0.0001267072,0.0016541051],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99908984,0.0003834657,0.00006732667,0.00027868524,0.000114020106,0.00006678789],"domain_scores_gemma":[0.9910782,0.0070818583,0.00048780913,0.00057558075,0.00060945423,0.00016708125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020717168,0.0014168385,0.0011383394,0.002611116,0.0003553595,0.0014253238,0.0016023039,0.001319573,0.0010706144],"category_scores_gemma":[0.011490652,0.00073292584,0.00092438515,0.0015612422,0.00035796931,0.0026617113,0.0004963086,0.001916262,0.0007120031],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006851533,0.002145999,0.092398874,0.00031600977,0.00074540934,0.00017993763,0.000729521,0.6129736,0.0040663523,0.0040604407,0.006188614,0.27551004],"study_design_scores_gemma":[0.000013306203,0.000039920265,0.0012771579,0.0000074823483,0.000023663773,0.000023168594,0.000018707397,0.9963897,0.00031697308,0.0016916641,0.00019031094,0.000007922449],"about_ca_topic_score_codex":0.01299429,"about_ca_topic_score_gemma":0.019692603,"teacher_disagreement_score":0.01299429,"about_ca_system_score_codex":0.001095235,"about_ca_system_score_gemma":0.000887898,"threshold_uncertainty_score":0.025837302},"labels":[],"label_agreement":null},{"id":"W1515236857","doi":"","title":"Poly-co: an unsupervised co-reference detection system","year":2010,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Conditional random field; Computer science; Artificial intelligence; Natural language processing; Named-entity recognition; Co-occurrence; Speech recognition; Pattern recognition (psychology); Engineering","score_opus":0.016266787439585633,"score_gpt":0.242484993553053,"score_spread":0.22621820611346738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1515236857","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016874854,0.0015838682,0.72836107,0.0004544548,0.00053785,0.00041332789,0.017279,0.2200139,0.0144816935],"genre_scores_gemma":[0.13488333,0.0006575336,0.7626226,0.0006488197,0.0004920422,0.0008332385,0.06618605,0.0082356045,0.025440812],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974069,0.00047140114,0.00014513634,0.0012132999,0.00059181673,0.00017147377],"domain_scores_gemma":[0.99553704,0.0008122108,0.0002782719,0.0020317165,0.0011249599,0.00021580842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018636581,0.0018836304,0.00207873,0.004261161,0.0015917675,0.0024621906,0.003765936,0.0017477592,0.0104507655],"category_scores_gemma":[0.0046559097,0.00091310777,0.0008091469,0.0038522205,0.00088502164,0.0050702807,0.0040633935,0.0017044992,0.014561428],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012075811,0.0005381131,0.004943156,0.0007892894,0.00032553409,0.0006376743,0.0005052661,0.0073007713,0.045946218,0.008944362,0.3337755,0.5950866],"study_design_scores_gemma":[0.00026977732,0.00039834564,0.008957168,0.00011210477,0.00030209834,0.0022877324,0.00022550982,0.54449713,0.10566555,0.016290411,0.32061934,0.00037479145],"about_ca_topic_score_codex":0.006725712,"about_ca_topic_score_gemma":0.013279954,"teacher_disagreement_score":0.0104507655,"about_ca_system_score_codex":0.00081349764,"about_ca_system_score_gemma":0.0020390733,"threshold_uncertainty_score":0.034961224},"labels":[],"label_agreement":null},{"id":"W1522404235","doi":"","title":"Near-synonym Lexical Choice in Latent Semantic Space","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Latent semantic analysis; Computer science; Artificial intelligence; Natural language processing; Synonym (taxonomy); Probabilistic latent semantic analysis; Curse of dimensionality; Context (archaeology); Space (punctuation); Semantic space; Task (project management); Representation (politics)","score_opus":0.022122818924413572,"score_gpt":0.2580109281140938,"score_spread":0.23588810918968026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1522404235","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10885506,0.0007614037,0.8862118,0.0009189815,0.000057610705,0.00012396321,0.00054442533,0.0008829086,0.0016438065],"genre_scores_gemma":[0.7992336,0.00041258827,0.19710687,0.00017301367,0.00011562758,0.00017169672,0.0008549602,0.00011329397,0.0018182084],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99379843,0.003041784,0.00035736462,0.0015142007,0.001002113,0.00028614525],"domain_scores_gemma":[0.98664933,0.009708812,0.0013223345,0.001547724,0.0004672815,0.0003045081],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005852315,0.00094183587,0.0023172274,0.0029679888,0.0015264956,0.0039755725,0.0022862877,0.0020579277,0.0029210197],"category_scores_gemma":[0.015857501,0.00068375806,0.0020941922,0.004708386,0.0023456048,0.011963487,0.004541908,0.0026701067,0.0009105622],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00144772,0.00094478624,0.030405989,0.0012784225,0.0006665716,0.0016267671,0.005384932,0.15424423,0.02414608,0.3433433,0.010079224,0.42643198],"study_design_scores_gemma":[0.00006624421,0.000111145,0.0019145644,0.00005251825,0.00008127301,0.0005304752,0.00050503714,0.7372873,0.004122598,0.25156546,0.0036778096,0.000085530715],"about_ca_topic_score_codex":0.0021279096,"about_ca_topic_score_gemma":0.0028877202,"teacher_disagreement_score":0.005852315,"about_ca_system_score_codex":0.001192972,"about_ca_system_score_gemma":0.0012186952,"threshold_uncertainty_score":0.030950367},"labels":[],"label_agreement":null},{"id":"W1530262931","doi":"10.1007/11765448_11","title":"Using Semantic Constraints to Improve Question Answering","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Redundancy (engineering); Question answering; Information retrieval; Ranking (information retrieval); Rank (graph theory); Artificial intelligence; Natural language processing; Semantic Web","score_opus":0.02340020090838834,"score_gpt":0.2666794486687881,"score_spread":0.24327924776039977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1530262931","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017872361,0.0022543161,0.9515891,0.0019219262,0.0005749865,0.00027676116,0.0028008483,0.009116707,0.0135929445],"genre_scores_gemma":[0.16605385,0.001370023,0.80915713,0.0006877706,0.0004791955,0.00041028473,0.012415063,0.0017655732,0.007661135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957343,0.0017832697,0.00032502817,0.0007654264,0.0011908946,0.00020110467],"domain_scores_gemma":[0.9856296,0.010359491,0.00035182785,0.0015791728,0.0017960347,0.0002838641],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003960041,0.001311461,0.0015161126,0.0021937129,0.0011269099,0.00282219,0.0027798251,0.0016957654,0.016430728],"category_scores_gemma":[0.02217713,0.000940644,0.0015574886,0.0029529724,0.0008505856,0.010785346,0.0033831738,0.0036386964,0.0042750207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004384602,0.00043547875,0.0013961171,0.0013408465,0.0001702893,0.00017001962,0.0011518225,0.018305216,0.018653518,0.10535361,0.08772308,0.7648615],"study_design_scores_gemma":[0.00017501399,0.00012973034,0.0011903347,0.00025466306,0.0002955301,0.0002869385,0.0005256254,0.5427093,0.022586513,0.31830388,0.11344532,0.000097132965],"about_ca_topic_score_codex":0.004457059,"about_ca_topic_score_gemma":0.00667999,"teacher_disagreement_score":0.016430728,"about_ca_system_score_codex":0.00094747776,"about_ca_system_score_gemma":0.0014902125,"threshold_uncertainty_score":0.05496627},"labels":[],"label_agreement":null},{"id":"W1531174292","doi":"","title":"Domain Adaptation to Summarize Human Conversations","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Domain adaptation; Leverage (statistics); Computer science; Domain (mathematical analysis); Adaptation (eye); Labeled data; Artificial intelligence; Natural language processing; Information retrieval; Machine learning; Classifier (UML); Psychology; Mathematics","score_opus":0.02784659043970881,"score_gpt":0.26229919776693444,"score_spread":0.23445260732722564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1531174292","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04980624,0.00077595626,0.94114804,0.00028120592,0.00014276232,0.0002607236,0.0005398899,0.0041896463,0.0028554874],"genre_scores_gemma":[0.4548468,0.00058028445,0.5334625,0.000320646,0.00030848462,0.000622114,0.005429286,0.00046208,0.003967749],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99746704,0.0014489415,0.000112025926,0.0006939423,0.00018284321,0.000095196556],"domain_scores_gemma":[0.99524754,0.0025139304,0.00038889665,0.00089008466,0.0008183212,0.00014124799],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027948536,0.0010319784,0.0007226614,0.0015045695,0.0006777101,0.000970196,0.00086708,0.00087944826,0.0018823365],"category_scores_gemma":[0.010724878,0.00031808385,0.00081788254,0.0013404124,0.00044560715,0.0021517528,0.0012144934,0.0015939024,0.001994423],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006524862,0.00073481083,0.0054282285,0.0005696196,0.00028833127,0.0001269549,0.001608557,0.08228777,0.06826302,0.005014725,0.012665959,0.8223597],"study_design_scores_gemma":[0.00007895903,0.00052738265,0.0069093853,0.00006759529,0.00018420887,0.00026271652,0.0009887868,0.89839464,0.053658966,0.019184638,0.019646408,0.000096275435],"about_ca_topic_score_codex":0.0009513104,"about_ca_topic_score_gemma":0.0016367635,"teacher_disagreement_score":0.0027948536,"about_ca_system_score_codex":0.00049075164,"about_ca_system_score_gemma":0.0006681058,"threshold_uncertainty_score":0.01478076},"labels":[],"label_agreement":null},{"id":"W1538369613","doi":"10.1007/978-3-540-72665-4_25","title":"Question Answering Summarization of Multiple Biomedical Documents","year":2007,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Automatic summarization; Information retrieval; Computer science; Question answering","score_opus":0.02123581115807662,"score_gpt":0.27439893159899154,"score_spread":0.25316312044091493,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1538369613","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11968552,0.021077776,0.817051,0.0051710806,0.0012889251,0.00094070233,0.0155565925,0.009855232,0.0093731955],"genre_scores_gemma":[0.29195875,0.006242176,0.6199294,0.0009764657,0.0023990422,0.00062226946,0.062957406,0.00073199807,0.014182421],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980332,0.00069591554,0.0002296195,0.00044683268,0.00045338972,0.00014102092],"domain_scores_gemma":[0.9950014,0.0031508985,0.0003178676,0.00035164642,0.0010382439,0.0001399033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027227197,0.0014925351,0.0018686871,0.0047633355,0.0008393907,0.0021390691,0.001602484,0.0018711665,0.00583272],"category_scores_gemma":[0.0069822627,0.0004356337,0.0014617752,0.003202547,0.00034863633,0.0022584274,0.0012655739,0.0012708955,0.0029610954],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094328204,0.00041093046,0.0019657896,0.0024867142,0.00052029773,0.0009514408,0.0011255764,0.012163469,0.072643094,0.005072319,0.0640952,0.83762187],"study_design_scores_gemma":[0.0004092861,0.0017963562,0.020410802,0.00089681253,0.0030179687,0.0041279844,0.0027744847,0.55012876,0.1577712,0.058725696,0.19967602,0.00026464034],"about_ca_topic_score_codex":0.001319465,"about_ca_topic_score_gemma":0.0019180094,"teacher_disagreement_score":0.00583272,"about_ca_system_score_codex":0.0006536028,"about_ca_system_score_gemma":0.0009745725,"threshold_uncertainty_score":0.019512415},"labels":[],"label_agreement":null},{"id":"W1541367332","doi":"10.1007/978-3-642-21043-3_15","title":"Exploiting Conversational Features to Detect High-Quality Blog Comments","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"CRFS; Computer science; Conditional random field; Quality (philosophy); Natural language processing; Artificial intelligence; Moderation; Social media; Binary number; Information retrieval; Machine learning; World Wide Web; Mathematics","score_opus":0.04797135374995147,"score_gpt":0.27760490326209586,"score_spread":0.22963354951214437,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1541367332","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9070769,0.003561719,0.06817628,0.00066400407,0.00077045953,0.00028953326,0.0061085112,0.0032210718,0.010131649],"genre_scores_gemma":[0.9643561,0.0005649018,0.024344295,0.00009033568,0.000681986,0.00014534133,0.0042343703,0.00020647951,0.005376248],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909925,0.00019554807,0.0000619813,0.00017000888,0.00033104597,0.00014217541],"domain_scores_gemma":[0.9945998,0.0032389557,0.00057239784,0.00023304905,0.0010087995,0.0003469761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001020322,0.0006147478,0.00056873146,0.0041358904,0.0007759202,0.0013764028,0.0004409729,0.0008470446,0.0022078918],"category_scores_gemma":[0.005670626,0.00030240664,0.00042728856,0.0020617368,0.00018836213,0.0011609697,0.0009341456,0.00088378746,0.0020935864],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032433125,0.00086378295,0.10346387,0.0012840566,0.00040546086,0.0014086668,0.0027739685,0.0035541605,0.20954812,0.0015504723,0.032311343,0.6395928],"study_design_scores_gemma":[0.00022664567,0.0014032793,0.40605348,0.0002472934,0.0007249881,0.0030189694,0.004717023,0.4601492,0.07530383,0.0060768714,0.04180372,0.00027467392],"about_ca_topic_score_codex":0.001447572,"about_ca_topic_score_gemma":0.0040951152,"teacher_disagreement_score":0.0041358904,"about_ca_system_score_codex":0.0002623356,"about_ca_system_score_gemma":0.0003590863,"threshold_uncertainty_score":0.0073862076},"labels":[],"label_agreement":null},{"id":"W1541437296","doi":"","title":"Unsupervised Modeling of Dialog Acts in Asynchronous Conversations","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Dialog box; Conversation; Asynchronous communication; Hidden Markov model; Artificial intelligence; Graph; Probabilistic logic; Natural language processing; Theoretical computer science; World Wide Web; Linguistics","score_opus":0.040179922661745804,"score_gpt":0.2484000574484378,"score_spread":0.208220134786692,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1541437296","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07449213,0.0003451264,0.92225397,0.00028860933,0.000037064514,0.00012389502,0.00045802337,0.0005514678,0.0014497768],"genre_scores_gemma":[0.8428702,0.00037784214,0.1518146,0.0001163418,0.00014650315,0.0004261576,0.0012978625,0.0001461848,0.0028043012],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977718,0.0011745897,0.00009137913,0.0006582668,0.00019766223,0.000106252446],"domain_scores_gemma":[0.9922725,0.0058464394,0.0006438819,0.00056891027,0.00043877706,0.00022960966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002487353,0.0007669532,0.0007792956,0.0013288257,0.00086901995,0.0015578431,0.001621532,0.0012864035,0.0012922038],"category_scores_gemma":[0.010284185,0.00065489224,0.0012328646,0.00080914964,0.000916853,0.0026358927,0.0014613075,0.0016193828,0.0005606763],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004273601,0.0002912017,0.009046005,0.00034910094,0.00025435546,0.0002821342,0.0020678213,0.8043924,0.011484474,0.06284226,0.0026252156,0.105937585],"study_design_scores_gemma":[0.000006614625,0.000015137432,0.00068963453,0.000006940693,0.000010840501,0.000027011678,0.000037807957,0.9831924,0.0005460484,0.01489308,0.0005648338,0.000009590102],"about_ca_topic_score_codex":0.0046925982,"about_ca_topic_score_gemma":0.0066683437,"teacher_disagreement_score":0.0046925982,"about_ca_system_score_codex":0.00097025756,"about_ca_system_score_gemma":0.0011273823,"threshold_uncertainty_score":0.013154507},"labels":[],"label_agreement":null},{"id":"W1546791496","doi":"10.1609/aaai.v25i1.7975","title":"Using Semantic Cues to Learn Syntax","year":2011,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Army Research Laboratory; Air Force Research Laboratory; Army Research Office; Advanced Research Projects Agency; Defense Advanced Research Projects Agency","keywords":"Computer science; Natural language processing; Artificial intelligence; Syntactic predicate; Intuition; Predicate (mathematical logic); Exploit; Syntax; Programming language; Cognitive science","score_opus":0.2479244645014024,"score_gpt":0.32588280797772073,"score_spread":0.07795834347631833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1546791496","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011311899,0.0001613979,0.97857976,0.0003308382,0.000058047524,0.00007642299,0.00087933434,0.0053714453,0.0032308314],"genre_scores_gemma":[0.305199,0.00036865484,0.68049073,0.00036809716,0.00010633598,0.00032349408,0.0075035905,0.0012632174,0.0043768706],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988714,0.00040530012,0.00005342219,0.0004066905,0.00019364934,0.00006950147],"domain_scores_gemma":[0.99622464,0.0024851125,0.0001985053,0.0006141106,0.0003724506,0.00010520622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011966309,0.0012873681,0.00066924776,0.0025675665,0.0008400528,0.001283072,0.0016842786,0.0011210729,0.006390004],"category_scores_gemma":[0.0077811256,0.0008763324,0.0011783099,0.0015860006,0.0008805181,0.004449567,0.0022576968,0.002671608,0.0035652097],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020073615,0.00031029127,0.0050665163,0.00049194344,0.00016613147,0.00042560897,0.0007049536,0.072430536,0.028660413,0.07343585,0.027470212,0.79063684],"study_design_scores_gemma":[0.000049257877,0.000056395314,0.0010966522,0.00006578583,0.00006525166,0.00015652328,0.000115182906,0.8317744,0.014512255,0.13564815,0.01639563,0.00006438705],"about_ca_topic_score_codex":0.0030448444,"about_ca_topic_score_gemma":0.0075828223,"teacher_disagreement_score":0.006390004,"about_ca_system_score_codex":0.00091384235,"about_ca_system_score_gemma":0.001705065,"threshold_uncertainty_score":0.02137667},"labels":[],"label_agreement":null},{"id":"W1547627874","doi":"10.1016/s0306-4573(99)00048-5","title":"Passage-based query refinement","year":2000,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of Waterloo","funders":"","keywords":"Computer science; Ranking (information retrieval); Information retrieval; Relevance (law); Query expansion; Task (project management); Set (abstract data type); Interface (matter); Construct (python library); Relevance feedback; Function (biology); Term (time); Artificial intelligence; Image retrieval","score_opus":0.01105825729091549,"score_gpt":0.22562178257574456,"score_spread":0.21456352528482905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1547627874","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016842777,0.00088783127,0.96207166,0.0009104645,0.00021027966,0.0013678038,0.0018048348,0.011230284,0.0046739983],"genre_scores_gemma":[0.21789351,0.0008819121,0.7550229,0.0006586126,0.00030993504,0.0008625974,0.01119002,0.0021333422,0.0110471705],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98047596,0.007862065,0.001834868,0.0025495754,0.006285157,0.0009923839],"domain_scores_gemma":[0.9639657,0.017044283,0.0007206028,0.0080410065,0.0096236775,0.00060467527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012877833,0.0020577211,0.003381073,0.0071059396,0.0026364343,0.0037454383,0.0073756045,0.0023150458,0.017674698],"category_scores_gemma":[0.056487896,0.0014106038,0.0043628835,0.005527304,0.001897897,0.0072535556,0.0055197654,0.0037552298,0.0070207487],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028463101,0.0012756295,0.007682725,0.0024809544,0.0008884444,0.0014479692,0.0058599203,0.04164004,0.07516948,0.07167495,0.076329626,0.712704],"study_design_scores_gemma":[0.0005758929,0.000681725,0.0036119162,0.00022361835,0.0014368321,0.0016544681,0.0021418128,0.73217356,0.08165651,0.08198151,0.09353027,0.00033183346],"about_ca_topic_score_codex":0.023046475,"about_ca_topic_score_gemma":0.01646297,"teacher_disagreement_score":0.023046475,"about_ca_system_score_codex":0.0018909681,"about_ca_system_score_gemma":0.0043618022,"threshold_uncertainty_score":0.06810534},"labels":[],"label_agreement":null},{"id":"W1554237613","doi":"","title":"Identifying synonyms among distributionally similar words","year":2003,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":208,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing","score_opus":0.02478478907119317,"score_gpt":0.25200049803264957,"score_spread":0.2272157089614564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1554237613","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1357226,0.0028010118,0.8460842,0.0009406368,0.00035108044,0.0006070579,0.0021586143,0.0021667127,0.00916815],"genre_scores_gemma":[0.62162715,0.0011123698,0.36660483,0.00020216613,0.00048930483,0.0006698402,0.00536615,0.000301754,0.0036264476],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945226,0.0015569392,0.00056930294,0.0015497094,0.0014817527,0.0003197006],"domain_scores_gemma":[0.99038553,0.0052511184,0.0010953281,0.0012183314,0.0017767489,0.0002730372],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002551797,0.0010870461,0.0014969441,0.011279402,0.001927463,0.003203575,0.0020309223,0.0017900177,0.003807204],"category_scores_gemma":[0.016825687,0.00063314685,0.001474869,0.010040329,0.0015512414,0.0061991266,0.0036507866,0.0015808232,0.0019813983],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009982623,0.00058439415,0.07054107,0.0018643237,0.0009673253,0.0011874731,0.004985346,0.010217703,0.042310447,0.11594359,0.017583309,0.7328167],"study_design_scores_gemma":[0.00039228503,0.00082416425,0.10599729,0.0006216101,0.0010511186,0.008605893,0.0066067562,0.37337556,0.035965264,0.3781706,0.08769995,0.00068943173],"about_ca_topic_score_codex":0.0015126674,"about_ca_topic_score_gemma":0.0029627697,"teacher_disagreement_score":0.011279402,"about_ca_system_score_codex":0.00071430334,"about_ca_system_score_gemma":0.0012935817,"threshold_uncertainty_score":0.013495326},"labels":[],"label_agreement":null},{"id":"W1558536082","doi":"10.1007/978-3-319-30671-1_27","title":"Multi-document Summarization Based on Atomic Semantic Events and Their Temporal Relationships","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Novelty; Natural language processing; Event (particle physics); Sentence; Information retrieval; Artificial intelligence; Domain (mathematical analysis); Precision and recall; Set (abstract data type); Multi-document summarization; Salient","score_opus":0.029557536198477907,"score_gpt":0.24569082159476685,"score_spread":0.21613328539628895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1558536082","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064659804,0.0078548705,0.907836,0.0005488747,0.0008814547,0.00045661916,0.0046062684,0.007234741,0.005921465],"genre_scores_gemma":[0.2808633,0.0027748952,0.6843346,0.00015327142,0.0010115992,0.00039069107,0.017067695,0.00072056137,0.012683385],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99889296,0.00021373946,0.00015577512,0.00033054265,0.0003150149,0.0000920885],"domain_scores_gemma":[0.99728715,0.001131476,0.00023628653,0.00028241577,0.0009581156,0.00010457659],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009869311,0.0015707539,0.0014397523,0.0040380047,0.0010410829,0.0019597637,0.0009415464,0.0009941089,0.0041541937],"category_scores_gemma":[0.0029214532,0.00053446257,0.0012642876,0.0037860095,0.0002641225,0.0028252716,0.00096520444,0.0008711503,0.0024538017],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00161252,0.00026063214,0.0023729785,0.0013676962,0.0003535736,0.0006189018,0.0005695118,0.010987735,0.08096284,0.0054385327,0.026169378,0.8692857],"study_design_scores_gemma":[0.00023543603,0.00087902305,0.0141050415,0.00023376655,0.002239747,0.0012001061,0.0010739805,0.7907615,0.1073178,0.021363022,0.060356513,0.00023412061],"about_ca_topic_score_codex":0.0016314619,"about_ca_topic_score_gemma":0.003465865,"teacher_disagreement_score":0.0041541937,"about_ca_system_score_codex":0.0003987502,"about_ca_system_score_gemma":0.0006754771,"threshold_uncertainty_score":0.013897181},"labels":[],"label_agreement":null},{"id":"W156106634","doi":"","title":"A Symbolic Summarizer with 2 Steps of Sentence Selection for TAC 2009.","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Automatic summarization; Selection (genetic algorithm); Computer science; Sentence; Task (project management); Natural language processing; Quality (philosophy); Artificial intelligence; Competition (biology)","score_opus":0.014362240522970603,"score_gpt":0.2371830052548323,"score_spread":0.2228207647318617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W156106634","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0701497,0.0016320342,0.634854,0.0006716267,0.0005640469,0.0014312528,0.025486164,0.25375536,0.011455784],"genre_scores_gemma":[0.247195,0.0003841936,0.6662877,0.00020424812,0.00023856398,0.0012434621,0.065497994,0.003360394,0.0155884195],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990447,0.00036380443,0.00010367042,0.00020656883,0.00023344094,0.00004790484],"domain_scores_gemma":[0.99848026,0.0006195232,0.00011252152,0.00024838388,0.00045211727,0.00008728929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017703284,0.00088219356,0.00088544626,0.0011250039,0.00042156386,0.00085965346,0.0010410412,0.0006250499,0.009737088],"category_scores_gemma":[0.0056902594,0.000359449,0.00057361636,0.00089242676,0.00014492401,0.0011576264,0.00060682016,0.0007191323,0.0057801506],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016056122,0.00041767006,0.0030275926,0.0018343269,0.0005390628,0.00046612375,0.0009854458,0.014233112,0.10004105,0.0036905122,0.18368587,0.6894736],"study_design_scores_gemma":[0.001273609,0.0027672478,0.015760928,0.00014145438,0.00080766395,0.0013398028,0.0006817863,0.5088974,0.19098102,0.007778177,0.269253,0.00031793347],"about_ca_topic_score_codex":0.0021160105,"about_ca_topic_score_gemma":0.005333794,"teacher_disagreement_score":0.009737088,"about_ca_system_score_codex":0.00046895954,"about_ca_system_score_gemma":0.00063618424,"threshold_uncertainty_score":0.03257376},"labels":[],"label_agreement":null},{"id":"W1578729058","doi":"10.18653/v1/s15-2016","title":"TrWP: Text Relatedness using Word and Phrase Relatedness","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Phrase; Natural language processing; Word (group theory); Artificial intelligence; Computer science; Semantics (computer science); Speech recognition; Linguistics","score_opus":0.07341750464105991,"score_gpt":0.2770615194956916,"score_spread":0.20364401485463168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1578729058","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04360288,0.0016221163,0.8094012,0.00042767564,0.00034636015,0.0019999205,0.027580574,0.10181163,0.013207667],"genre_scores_gemma":[0.22854117,0.000607029,0.67709947,0.00027409318,0.00032442238,0.00232571,0.070679724,0.005609907,0.014538355],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960699,0.0010375227,0.00030859522,0.001285024,0.001081091,0.00021773929],"domain_scores_gemma":[0.9970503,0.0009807319,0.00028825956,0.00090506213,0.0006034901,0.00017227743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024573174,0.0027654187,0.0011932151,0.0076069064,0.0012014661,0.0019760346,0.0022287471,0.0017275336,0.011136492],"category_scores_gemma":[0.012032514,0.00080114487,0.0019815755,0.005100774,0.0005154371,0.006175933,0.004606425,0.0017746,0.0111301085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074841053,0.0006687951,0.01073439,0.001844122,0.0008160501,0.00037850076,0.0010670145,0.011576679,0.025394084,0.010267556,0.15589075,0.7806137],"study_design_scores_gemma":[0.00042624172,0.0014401183,0.04220759,0.00029909617,0.00070993387,0.0018504488,0.0010997002,0.6610106,0.050113246,0.07558897,0.16484512,0.00040894392],"about_ca_topic_score_codex":0.0032169016,"about_ca_topic_score_gemma":0.004952051,"teacher_disagreement_score":0.011136492,"about_ca_system_score_codex":0.0007423278,"about_ca_system_score_gemma":0.0011519459,"threshold_uncertainty_score":0.037255287},"labels":[],"label_agreement":null},{"id":"W1588617094","doi":"10.1007/978-3-642-15760-8_27","title":"Coverage-Based Methods for Distributional Stopword Selection in Text Segmentation","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Computer science; Selection (genetic algorithm); Segmentation; Measure (data warehouse); Representation (politics); Word (group theory); Artificial intelligence; Natural language processing; Information retrieval; Data mining; Mathematics","score_opus":0.024365051874352766,"score_gpt":0.31662918369622595,"score_spread":0.29226413182187316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1588617094","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022495681,0.0018387784,0.9679394,0.00018032198,0.00013542226,0.00014813346,0.00069262006,0.005327701,0.0012418845],"genre_scores_gemma":[0.28240082,0.0011981353,0.6954201,0.00030141766,0.000706482,0.00073314836,0.009569204,0.0038390823,0.005831536],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949956,0.0017128767,0.0005421738,0.0010389856,0.0013231175,0.00038725027],"domain_scores_gemma":[0.97731555,0.017281674,0.00067447295,0.0014710427,0.0028321,0.0004251769],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050143483,0.0019902675,0.0039537563,0.0077508707,0.0018922925,0.0029517596,0.0033343516,0.0031466342,0.0055112573],"category_scores_gemma":[0.017302478,0.0012300378,0.0019243931,0.007033762,0.0012733708,0.0040393933,0.003456311,0.002921167,0.0045527173],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012022702,0.0003023549,0.004019498,0.00072360516,0.0004654418,0.000325033,0.00063264405,0.040960487,0.025701055,0.008029514,0.014660847,0.9029772],"study_design_scores_gemma":[0.00008636521,0.00012794421,0.0022035681,0.000057584224,0.00016372133,0.00028199164,0.00020042021,0.96624535,0.009802453,0.016193934,0.0045740847,0.000062572326],"about_ca_topic_score_codex":0.0054921205,"about_ca_topic_score_gemma":0.011418797,"teacher_disagreement_score":0.0077508707,"about_ca_system_score_codex":0.000800426,"about_ca_system_score_gemma":0.0017380714,"threshold_uncertainty_score":0.026518703},"labels":[],"label_agreement":null},{"id":"W1596986901","doi":"","title":"Joint learning of words and meaning representations for open-text semantic parsing","year":2012,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":333,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Natural language processing; WordNet; Artificial intelligence; Parsing; Task (project management); Meaning (existential); Context (archaeology); Representation (politics); Natural language; Process (computing); Programming language","score_opus":0.10819543704019396,"score_gpt":0.337927349450232,"score_spread":0.22973191241003804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1596986901","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015939081,0.00025668106,0.97751087,0.00042768207,0.000048958547,0.00006118214,0.00033268862,0.0043561873,0.0010666895],"genre_scores_gemma":[0.38697553,0.00039034145,0.605617,0.00025089865,0.00012833864,0.00018812528,0.0032000213,0.0004991972,0.002750504],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99873847,0.00057448866,0.00007018543,0.0003884111,0.00014484016,0.00008355999],"domain_scores_gemma":[0.9975241,0.0013717648,0.000189382,0.0005873692,0.0002211015,0.00010623813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021733865,0.0010520216,0.0008713416,0.0016056433,0.0007688122,0.0017026333,0.0019173167,0.001648406,0.0026368706],"category_scores_gemma":[0.006784994,0.0006849588,0.0012619831,0.001926081,0.0011330966,0.0068956832,0.002863141,0.0030024895,0.0020667955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034636955,0.00045762656,0.005027874,0.00037239387,0.00022202304,0.00039703358,0.0013030875,0.123597816,0.018615087,0.07737226,0.022370307,0.7499181],"study_design_scores_gemma":[0.000020491107,0.000042646327,0.0007089802,0.000025396117,0.000035821915,0.00011703155,0.0001553111,0.87273914,0.007275514,0.113915086,0.0049326285,0.000031942207],"about_ca_topic_score_codex":0.0018895711,"about_ca_topic_score_gemma":0.0032245682,"teacher_disagreement_score":0.0026368706,"about_ca_system_score_codex":0.0008147288,"about_ca_system_score_gemma":0.0010679745,"threshold_uncertainty_score":0.0114940405},"labels":[],"label_agreement":null},{"id":"W1602711325","doi":"10.48550/arxiv.1405.0603","title":"Extracting Family Relationship Networks from Novels","year":2014,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science","score_opus":0.13805781632773537,"score_gpt":0.19570764962359652,"score_spread":0.05764983329586115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1602711325","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43110913,0.009480645,0.46883553,0.0020139744,0.00038610303,0.0006812229,0.047737535,0.0040043457,0.035751496],"genre_scores_gemma":[0.7699226,0.0040122108,0.17727618,0.00011989448,0.00045479805,0.00055209297,0.03841918,0.00034987414,0.008893197],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992198,0.0001769696,0.00007071028,0.00029239221,0.00018239839,0.000057775265],"domain_scores_gemma":[0.9963744,0.002249254,0.00051934645,0.0003423598,0.00040024697,0.00011440735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00090541114,0.0008053477,0.00037850617,0.008788816,0.0011080776,0.0013158652,0.0007326133,0.0008206927,0.0030594117],"category_scores_gemma":[0.0062235803,0.00042928275,0.0006973625,0.0063220607,0.00038061588,0.0026391388,0.00095773133,0.00094205077,0.0017342938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005066037,0.0004019783,0.091078445,0.0016814106,0.000368478,0.0041056215,0.010352159,0.023557788,0.030668832,0.032740954,0.05192252,0.7526152],"study_design_scores_gemma":[0.00006052406,0.0002027617,0.17071103,0.0006323458,0.00056865363,0.006548456,0.009207375,0.3578997,0.024856864,0.07776499,0.35137725,0.00017009053],"about_ca_topic_score_codex":0.0042076167,"about_ca_topic_score_gemma":0.009213669,"teacher_disagreement_score":0.008788816,"about_ca_system_score_codex":0.00069848495,"about_ca_system_score_gemma":0.000599156,"threshold_uncertainty_score":0.010234773},"labels":[],"label_agreement":null},{"id":"W1605562353","doi":"10.1007/11510888_46","title":"An Approach to Mining Picture Objects Based on Textual Cues","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Semantics (computer science); Domain (mathematical analysis); Information retrieval; Task (project management); Natural language processing; Process (computing); Text processing; Artificial intelligence; Meaning (existential); Programming language","score_opus":0.02463087384928613,"score_gpt":0.2521876456404516,"score_spread":0.22755677179116546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1605562353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031551093,0.0034444032,0.9470038,0.00067579263,0.00020256935,0.00077088305,0.0038480323,0.0075269714,0.00497641],"genre_scores_gemma":[0.081223875,0.0011947502,0.9019855,0.0002506638,0.00021515955,0.00060018373,0.0074045896,0.000422313,0.006702903],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985923,0.00013663995,0.000110437555,0.000510849,0.0005097693,0.00014000053],"domain_scores_gemma":[0.9984518,0.00056236575,0.00012870059,0.00022195454,0.0005330738,0.00010209933],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009207662,0.0016549126,0.001785762,0.00823132,0.0013469722,0.0026552824,0.0035832892,0.0023918597,0.0051894668],"category_scores_gemma":[0.0029566824,0.00088166416,0.0021582465,0.008665461,0.00088777655,0.004167942,0.0027000743,0.0015706829,0.0049480232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043626645,0.00033659147,0.0035956202,0.0006512892,0.00016863515,0.00031783117,0.00056102726,0.0024383885,0.05220749,0.0052136285,0.015756585,0.91831666],"study_design_scores_gemma":[0.00030774015,0.0007709454,0.019184835,0.00035486813,0.0009860935,0.0030595558,0.0039106514,0.7245668,0.09712819,0.05429109,0.09512318,0.00031599836],"about_ca_topic_score_codex":0.008137692,"about_ca_topic_score_gemma":0.015079631,"teacher_disagreement_score":0.00823132,"about_ca_system_score_codex":0.0008247167,"about_ca_system_score_gemma":0.0013590506,"threshold_uncertainty_score":0.017360508},"labels":[],"label_agreement":null},{"id":"W161156596","doi":"","title":"Generation of Formal and Informal Sentences","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Formality; Computer science; Style (visual arts); Task (project management); Natural language processing; Natural language generation; Set (abstract data type); Quality (philosophy); Formal language; Artificial intelligence; Natural language; Linguistics; Programming language","score_opus":0.08998097047917206,"score_gpt":0.24044867653227348,"score_spread":0.15046770605310142,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W161156596","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.063158154,0.00024524584,0.9183172,0.00039243695,0.0002650994,0.0013528381,0.002986011,0.0071446346,0.006138448],"genre_scores_gemma":[0.20587288,0.00020941767,0.7795359,0.00022494809,0.00013504476,0.001275422,0.007376404,0.0016020606,0.0037678706],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99532074,0.0023102702,0.00047598523,0.0007316507,0.0010280324,0.00013335458],"domain_scores_gemma":[0.97711474,0.015397964,0.00132042,0.0027345808,0.0031114072,0.00032088652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039578034,0.0015495016,0.00079411454,0.0013783455,0.0006982237,0.0010815581,0.0011156624,0.0008425195,0.0066170255],"category_scores_gemma":[0.027104445,0.000520454,0.0010899745,0.00078340963,0.0005786794,0.0018041786,0.001778042,0.0009877844,0.0024970495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000876243,0.00064002175,0.009531114,0.003538909,0.00020748128,0.0020270252,0.0070401826,0.028679527,0.10879344,0.05009633,0.0385855,0.7499842],"study_design_scores_gemma":[0.0006531713,0.0013462116,0.009121268,0.000725912,0.0004118993,0.003987868,0.002681011,0.46861303,0.19761819,0.105162114,0.20924369,0.00043568364],"about_ca_topic_score_codex":0.00032329035,"about_ca_topic_score_gemma":0.00045398582,"teacher_disagreement_score":0.0066170255,"about_ca_system_score_codex":0.00058307836,"about_ca_system_score_gemma":0.00087426644,"threshold_uncertainty_score":0.022136211},"labels":[],"label_agreement":null},{"id":"W161376815","doi":"","title":"Identifying Relationships Between Entities in Text for Complex Interactive Question Answering Task.","year":2006,"lang":"en","type":"article","venue":"Bilkent University Institutional Repository (Bilkent University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Section (typography); Task (project management); Question answering; Natural language processing; Artificial intelligence; Information retrieval; Linguistics","score_opus":0.04670287172852092,"score_gpt":0.23165165218117137,"score_spread":0.18494878045265045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W161376815","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38255623,0.0022674887,0.5943519,0.0019206735,0.00011337654,0.0009068597,0.004924319,0.0025033564,0.01045589],"genre_scores_gemma":[0.77504855,0.00051142904,0.21033792,0.00015875623,0.000098958204,0.00037505224,0.009593681,0.00021728728,0.0036582758],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99583495,0.0022319406,0.00023621105,0.0010253078,0.0005203956,0.00015114047],"domain_scores_gemma":[0.97994393,0.0153370155,0.0018183434,0.0014454902,0.0010406392,0.00041452373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004763711,0.0007306729,0.00065524277,0.0035923442,0.0014716723,0.0021722307,0.0012710927,0.0020993084,0.008114694],"category_scores_gemma":[0.027333468,0.0003877494,0.0007490623,0.0027132262,0.0006048679,0.0075615877,0.0024197532,0.0012483466,0.003041241],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014006696,0.0009518174,0.0922871,0.002682531,0.00043992684,0.001937282,0.02744004,0.013885255,0.19290927,0.051171154,0.02127899,0.59361595],"study_design_scores_gemma":[0.00018201902,0.0013289856,0.12863976,0.0005924079,0.0008985113,0.0040293247,0.019134993,0.44251475,0.08351906,0.14190106,0.17697461,0.00028446876],"about_ca_topic_score_codex":0.0009947086,"about_ca_topic_score_gemma":0.00148362,"teacher_disagreement_score":0.008114694,"about_ca_system_score_codex":0.0005768135,"about_ca_system_score_gemma":0.0005818371,"threshold_uncertainty_score":0.02714634},"labels":[],"label_agreement":null},{"id":"W1633230169","doi":"10.1609/aimag.v35i1.2502","title":"Natural Language Access to Enterprise Data","year":2014,"lang":"en","type":"article","venue":"AI Magazine","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Siemens (Canada)","funders":"","keywords":"Computer science; Syntax; Data control language; Semantics (computer science); Natural language; Set (abstract data type); Data access; Enterprise data management; Question answering; Interpretation (philosophy); Query language; Data manipulation language; Programming language; Information retrieval; Enterprise information system; Database; World Wide Web; Data science; Artificial intelligence; Web search query; Query by Example","score_opus":0.026276509443169737,"score_gpt":0.3160102430337011,"score_spread":0.2897337335905314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1633230169","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035862975,0.002903619,0.87635404,0.0069953827,0.00018801648,0.00043013028,0.012779398,0.03278613,0.031700246],"genre_scores_gemma":[0.41900006,0.003325294,0.5102888,0.0037802532,0.0005647548,0.00062975206,0.04604773,0.0027305295,0.013632904],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954464,0.0019119676,0.00042063367,0.0008571901,0.0011535926,0.00021023552],"domain_scores_gemma":[0.98889124,0.00722156,0.00044952935,0.0021475204,0.0010725858,0.00021750592],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043046866,0.00048807552,0.000694237,0.0032602376,0.0010854645,0.0033598228,0.0016789832,0.001092792,0.0070826835],"category_scores_gemma":[0.018645732,0.0004677748,0.0008317572,0.0032009615,0.001058281,0.0069751977,0.00398255,0.0013312332,0.0024048027],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043177782,0.00034614146,0.0062952014,0.0016900981,0.00019329829,0.0017420371,0.009607168,0.012178221,0.025989894,0.23511729,0.1465044,0.55990446],"study_design_scores_gemma":[0.00007103181,0.00007073404,0.0022748006,0.00023775826,0.000067618144,0.00086152885,0.0018120172,0.11525538,0.0134700965,0.28405854,0.58171487,0.000105504434],"about_ca_topic_score_codex":0.0049111675,"about_ca_topic_score_gemma":0.0054673585,"teacher_disagreement_score":0.0070826835,"about_ca_system_score_codex":0.0009779414,"about_ca_system_score_gemma":0.0014222974,"threshold_uncertainty_score":0.02369392},"labels":[],"label_agreement":null},{"id":"W1633328346","doi":"10.1007/978-3-642-21043-3_26","title":"Comparison of Semantic Similarity for Different Languages Using the Google n-gram Corpus and Second-Order Co-occurrence Measures","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Natural language processing; Similarity (geometry); German; Semantic similarity; Artificial intelligence; Word (group theory); n-gram; Word order; Language model; Linguistics","score_opus":0.07990233339630563,"score_gpt":0.32741458524636763,"score_spread":0.24751225185006198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1633328346","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.912475,0.0033053008,0.021828847,0.00033589322,0.00038090663,0.00023869988,0.041627087,0.0038724965,0.015935743],"genre_scores_gemma":[0.815459,0.0013332709,0.05328684,0.00007514988,0.00010912919,0.0003507056,0.12556483,0.000872025,0.0029489521],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978199,0.00067339826,0.00030381806,0.00033904376,0.0007038592,0.00016005475],"domain_scores_gemma":[0.9939115,0.0035022984,0.00024203544,0.000492137,0.0015957454,0.00025628821],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014060878,0.001016937,0.000928979,0.011105275,0.0014125804,0.0023443364,0.0006835499,0.0008859193,0.0033713374],"category_scores_gemma":[0.00996111,0.00024952926,0.0010954103,0.010622186,0.00063937664,0.0032782648,0.0015177758,0.00094159046,0.0023794125],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0077606137,0.0013637854,0.09230053,0.008341644,0.0016044788,0.0024780247,0.0068240273,0.014926912,0.08588388,0.016002258,0.09216452,0.6703493],"study_design_scores_gemma":[0.0005808832,0.0019420268,0.39580077,0.0011365766,0.0018032795,0.0068636057,0.018475277,0.30768263,0.09425928,0.025739688,0.14483136,0.00088460644],"about_ca_topic_score_codex":0.008607333,"about_ca_topic_score_gemma":0.015072303,"teacher_disagreement_score":0.011105275,"about_ca_system_score_codex":0.00075162714,"about_ca_system_score_gemma":0.0012664667,"threshold_uncertainty_score":0.01711446},"labels":[],"label_agreement":null},{"id":"W1649582405","doi":"","title":"Interactive Natural Language Query Construction for Report Generation","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Response Biomedical (Canada); Simon Fraser University","funders":"","keywords":"Computer science; Question answering; Natural language user interface; Natural language; Interface (matter); Natural language generation; Domain (mathematical analysis); Grammar; Point (geometry); Natural language processing; User interface; Information retrieval; World Wide Web; Selection (genetic algorithm); Artificial intelligence; Linguistics; Programming language","score_opus":0.02300953650142089,"score_gpt":0.29442933638425084,"score_spread":0.27141979988282994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1649582405","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01261389,0.0001739035,0.94043416,0.00052191,0.00004669412,0.00052652915,0.00096809183,0.041607454,0.0031073743],"genre_scores_gemma":[0.183246,0.00020255707,0.80470294,0.0002656934,0.00007727261,0.00076580804,0.0053841593,0.0030313623,0.002324255],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98784965,0.007552531,0.0008002745,0.0012284756,0.0021767186,0.00039227307],"domain_scores_gemma":[0.9683405,0.024156792,0.0008383992,0.0033754543,0.0027218617,0.0005669367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010470084,0.0014474115,0.0009921923,0.0020082204,0.0009627572,0.0025941408,0.0031880082,0.0017123414,0.012652256],"category_scores_gemma":[0.033905648,0.0007634556,0.0016600791,0.0013027262,0.0015248001,0.0043913126,0.0035054635,0.001532854,0.0050088037],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012519187,0.001276839,0.009122455,0.0019437235,0.00021713818,0.0015268694,0.01117323,0.033340085,0.06916239,0.113660075,0.09462614,0.66269916],"study_design_scores_gemma":[0.00041138934,0.0005119257,0.0025915166,0.00023784062,0.00015823428,0.0011439939,0.002191077,0.62238294,0.08274842,0.105770946,0.18159083,0.0002608887],"about_ca_topic_score_codex":0.0027900855,"about_ca_topic_score_gemma":0.0026786535,"teacher_disagreement_score":0.012652256,"about_ca_system_score_codex":0.0013631807,"about_ca_system_score_gemma":0.0016036654,"threshold_uncertainty_score":0.0553717},"labels":[],"label_agreement":null},{"id":"W1654173042","doi":"","title":"Unsupervised Modeling of Twitter Conversations","year":2010,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":434,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Microsoft Research","keywords":"Computer science; Conversation; Task (project management); Cluster analysis; Visualization; Domain (mathematical analysis); Artificial intelligence; Social media; Natural language processing; Microblogging; Topic model; Data science; World Wide Web; Linguistics","score_opus":0.03234468296183593,"score_gpt":0.24750362760771497,"score_spread":0.21515894464587904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1654173042","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.110560045,0.00029920353,0.88134146,0.000597939,0.000055592176,0.00015435116,0.0016932891,0.0010687541,0.0042293062],"genre_scores_gemma":[0.88486886,0.00030926953,0.106567286,0.00009910824,0.00015732228,0.00042297217,0.0029327273,0.00025983166,0.0043827035],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989691,0.00049458205,0.00004307236,0.0002786739,0.00012326211,0.00009131733],"domain_scores_gemma":[0.99767977,0.0014039453,0.00026836884,0.00028774454,0.00027171,0.000088542045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010084535,0.00056476647,0.0005435738,0.0011609265,0.0005663057,0.0011835509,0.0011183013,0.00070922857,0.0014287919],"category_scores_gemma":[0.006810176,0.00041084836,0.00079664344,0.0008705125,0.0006282655,0.0019167637,0.0010344357,0.0010797118,0.00068796164],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054337684,0.00023882589,0.026982067,0.00035844522,0.00026298471,0.00044979478,0.0034338639,0.6984304,0.019616855,0.09021916,0.009850348,0.1496139],"study_design_scores_gemma":[0.000004942787,0.000009051157,0.00121353,0.000005403518,0.000006131157,0.000025954925,0.00006912276,0.9860108,0.0006517785,0.010674609,0.0013205433,0.00000821147],"about_ca_topic_score_codex":0.0058417637,"about_ca_topic_score_gemma":0.007299322,"teacher_disagreement_score":0.0058417637,"about_ca_system_score_codex":0.0007663103,"about_ca_system_score_gemma":0.00076173764,"threshold_uncertainty_score":0.011615515},"labels":[],"label_agreement":null},{"id":"W1660606671","doi":"","title":"Summarize What You Are Interested In: An Optimization Framework for Interactive Personalized Summarization","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Automatic summarization; Personalization; Computer science; Consistency (knowledge bases); Preference; Information retrieval; Multi-document summarization; Personalized search; World Wide Web; Artificial intelligence","score_opus":0.07522027830310149,"score_gpt":0.29889772329250847,"score_spread":0.22367744498940698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1660606671","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053973207,0.00043599049,0.99173427,0.00024830914,0.00002129859,0.000075730015,0.000131363,0.000798896,0.001156901],"genre_scores_gemma":[0.2182679,0.00068422465,0.7731851,0.0002136928,0.00020402712,0.0004264148,0.0009757554,0.00047563633,0.0055673164],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991536,0.00034204934,0.000060034425,0.00018651091,0.00019401335,0.00006378429],"domain_scores_gemma":[0.99893814,0.0005877392,0.00011933167,0.000094232724,0.00021018155,0.000050358394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021726503,0.0012322582,0.0011464412,0.0011372984,0.0004984862,0.0013497068,0.001274608,0.0010814756,0.0032560145],"category_scores_gemma":[0.0041340096,0.0005090277,0.00075976213,0.001223097,0.00043452528,0.0018312983,0.0009763885,0.0010457619,0.0008958115],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019657821,0.00013956669,0.0006206795,0.00028273553,0.000117475494,0.00014515422,0.00036875662,0.6634222,0.0075660087,0.024678664,0.011258091,0.2912041],"study_design_scores_gemma":[0.000023161794,0.000070164104,0.0001655328,0.000012906204,0.000030450366,0.0000426793,0.000033547065,0.9864811,0.0012567998,0.0086858105,0.0031854487,0.000012323023],"about_ca_topic_score_codex":0.0035481432,"about_ca_topic_score_gemma":0.0052244114,"teacher_disagreement_score":0.0035481432,"about_ca_system_score_codex":0.00086337584,"about_ca_system_score_gemma":0.0010645838,"threshold_uncertainty_score":0.011490226},"labels":[],"label_agreement":null},{"id":"W1662133657","doi":"10.1613/jair.2934","title":"From Frequency to Meaning: Vector Space Models of Semantics","year":2010,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":2883,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Semantics (computer science); Meaning (existential); Computer science; Context (archaeology); Space (punctuation); Perspective (graphical); Artificial intelligence; Programming language; Psychology","score_opus":0.2081680926678031,"score_gpt":0.4142808458992752,"score_spread":0.20611275323147213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1662133657","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03187528,0.005143562,0.9355267,0.0052976403,0.00025376538,0.00011212847,0.0010666334,0.00047443938,0.020249868],"genre_scores_gemma":[0.7229849,0.0054058842,0.258254,0.0008200248,0.0010487897,0.0005380917,0.0014796007,0.00023994483,0.009228783],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975967,0.0013430915,0.00014431299,0.00035091175,0.00043023215,0.00013462025],"domain_scores_gemma":[0.9927769,0.00515273,0.00053813757,0.00070886395,0.00062103686,0.00020228051],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028091045,0.000881384,0.00095164933,0.0044672242,0.00092080265,0.0058356994,0.001745973,0.0013750126,0.0053647794],"category_scores_gemma":[0.018301178,0.00041647107,0.0015363912,0.005264552,0.0038302012,0.01662653,0.002092414,0.0022276328,0.0010126512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006323949,0.000025757972,0.0011666202,0.00009833342,0.0000396696,0.000048510527,0.0009852167,0.010375412,0.00017360049,0.9436175,0.0022763766,0.04112974],"study_design_scores_gemma":[0.000007608036,0.000013824832,0.00019851633,0.000024517338,0.000008953666,0.000045042925,0.00014986271,0.0341371,0.000044009954,0.9621933,0.0031675901,0.000009558034],"about_ca_topic_score_codex":0.004393936,"about_ca_topic_score_gemma":0.0028052628,"teacher_disagreement_score":0.0058356994,"about_ca_system_score_codex":0.0018440483,"about_ca_system_score_gemma":0.0011461966,"threshold_uncertainty_score":0.017946959},"labels":[],"label_agreement":null},{"id":"W166397155","doi":"10.1609/icwsm.v5i1.14135","title":"Extracting Meta Statements from the Blogosphere","year":2021,"lang":"en","type":"article","venue":"Proceedings of the International AAAI Conference on Web and Social Media","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Conditional random field; Computer science; Blogosphere; Information extraction; Relationship extraction; Statement (logic); Information retrieval; Classifier (UML); Precision and recall; Natural language processing; Artificial intelligence; Context (archaeology); Relation (database); Metadata; World Wide Web; Data mining; The Internet; Linguistics","score_opus":0.08824400931823363,"score_gpt":0.3052413327400337,"score_spread":0.21699732342180006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W166397155","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37068424,0.004172677,0.52657276,0.0025280935,0.00058219477,0.0006956114,0.044197544,0.026454996,0.024111928],"genre_scores_gemma":[0.66460824,0.002461676,0.27053577,0.00025364046,0.0006665504,0.00033124504,0.0542569,0.0007324346,0.006153635],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918646,0.00015186738,0.00009142993,0.00018028404,0.00029206974,0.00009784772],"domain_scores_gemma":[0.99541783,0.0026550433,0.00058069703,0.00041682846,0.0008044216,0.00012515424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014032399,0.0014384483,0.0006110645,0.008618419,0.0008831443,0.0014582139,0.00054845336,0.0007828533,0.0019285007],"category_scores_gemma":[0.0045802905,0.00047794633,0.0008293138,0.005129169,0.00039244466,0.004138171,0.0010933658,0.00084297074,0.0016228358],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087812106,0.00037274582,0.042884048,0.0014219778,0.00018053179,0.0023052162,0.002577686,0.008861168,0.06254093,0.014402978,0.045018937,0.81855565],"study_design_scores_gemma":[0.0001547593,0.00048016803,0.09096317,0.00058625423,0.00066645467,0.0032620113,0.0036971576,0.51342833,0.12204548,0.08667924,0.17776394,0.00027307562],"about_ca_topic_score_codex":0.0025716238,"about_ca_topic_score_gemma":0.0046476373,"teacher_disagreement_score":0.008618419,"about_ca_system_score_codex":0.0005284245,"about_ca_system_score_gemma":0.0012491114,"threshold_uncertainty_score":0.007421136},"labels":[],"label_agreement":null},{"id":"W169330724","doi":"10.5591/978-1-57735-516-8/ijcai11-461","title":"Analysis of adjective-noun word pair extraction methods for online review summarization","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Automatic summarization; Adjective; Computer science; Word (group theory); Natural language processing; Noun; Artificial intelligence; Identification (biology); Information retrieval; Linguistics","score_opus":0.1312730152729287,"score_gpt":0.4087756528637743,"score_spread":0.2775026375908456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W169330724","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36349368,0.007186064,0.6143942,0.0005944049,0.0003644626,0.0013914237,0.0026583318,0.0064885304,0.0034288925],"genre_scores_gemma":[0.58302104,0.0010247634,0.4075792,0.000099765144,0.0002151266,0.0010278709,0.004987585,0.00035450558,0.001690127],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99067605,0.005515609,0.0010410668,0.00071654835,0.0018651875,0.00018554175],"domain_scores_gemma":[0.93163913,0.053439002,0.002548502,0.0017489593,0.010260028,0.00036441727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008952581,0.0011725593,0.0010589437,0.0053476472,0.0006227363,0.0013990226,0.0006764737,0.00066212454,0.0014182985],"category_scores_gemma":[0.050240856,0.00028633958,0.00097363093,0.0035712437,0.00023647356,0.0020004695,0.0006394647,0.00070847606,0.0011599322],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014454465,0.00032619757,0.021590553,0.0020640637,0.0006661764,0.00034357276,0.0012766427,0.008963502,0.051526062,0.0013905589,0.008080878,0.9023263],"study_design_scores_gemma":[0.00036841066,0.0022634945,0.085395284,0.00027675432,0.0014267568,0.0015645183,0.0020758833,0.77238846,0.10835317,0.0047083744,0.020874737,0.00030425968],"about_ca_topic_score_codex":0.001792487,"about_ca_topic_score_gemma":0.0028717467,"teacher_disagreement_score":0.008952581,"about_ca_system_score_codex":0.00056929956,"about_ca_system_score_gemma":0.00073938764,"threshold_uncertainty_score":0.047346354},"labels":[],"label_agreement":null},{"id":"W1757017496","doi":"10.1007/978-3-642-01307-2_100","title":"Boosting Biomedical Information Retrieval Performance through Citation Graph: An Empirical Study","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Information retrieval; PageRank; Linkage (software); Graph; Ranking (information retrieval); Boosting (machine learning); Hyperlink; Divergence-from-randomness model; Data mining; Citation; Artificial intelligence; Theoretical computer science; Web page; Probabilistic logic; World Wide Web","score_opus":0.03868304289903317,"score_gpt":0.29826593550915986,"score_spread":0.2595828926101267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1757017496","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9840757,0.0040824367,0.007701303,0.00046473314,0.0000973509,0.00007873858,0.00036024628,0.0004188378,0.0027207162],"genre_scores_gemma":[0.9928939,0.00067086733,0.004309663,0.000082847415,0.0001772051,0.000018008315,0.00061836507,0.000071201386,0.0011579455],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9972429,0.0016356807,0.00011250797,0.00029235997,0.00055389403,0.00016273028],"domain_scores_gemma":[0.9285571,0.061606634,0.0023986883,0.0028981038,0.0034734819,0.001065891],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0076235617,0.00067775697,0.001138802,0.0027307617,0.00062033336,0.0015863917,0.0008613532,0.0013544817,0.0021160024],"category_scores_gemma":[0.05121833,0.00019770481,0.0007758429,0.0031425045,0.00057933445,0.0022718539,0.00054561044,0.0009349073,0.001407142],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007546391,0.0063992115,0.34825844,0.0009973617,0.0014100934,0.0002513618,0.0005155701,0.06149741,0.012238019,0.0020419033,0.013115731,0.5457285],"study_design_scores_gemma":[0.0006375945,0.009187053,0.3780763,0.00014254444,0.0033087258,0.001323804,0.0004971372,0.565262,0.019755168,0.013306533,0.008316217,0.00018696734],"about_ca_topic_score_codex":0.0016401901,"about_ca_topic_score_gemma":0.0012690423,"teacher_disagreement_score":0.9972692,"about_ca_system_score_codex":0.0005212176,"about_ca_system_score_gemma":0.00067050033,"threshold_uncertainty_score":0.040317774},"labels":[],"label_agreement":null},{"id":"W1764253967","doi":"","title":"Fuzzy set theory-based belief processing for natural language texts","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Fuzzy logic; Artificial intelligence; Newspaper; Set (abstract data type); Fuzzy set; Information extraction; Natural language; Natural language processing; Natural (archaeology); Machine learning; Programming language","score_opus":0.015259565497885349,"score_gpt":0.2895544110047849,"score_spread":0.27429484550689953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1764253967","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013481248,0.0007752131,0.9819264,0.00067707495,0.00004540448,0.00012474704,0.00023431933,0.0003002139,0.0024354286],"genre_scores_gemma":[0.58403707,0.0010397122,0.40940025,0.0002448026,0.00016584594,0.00046615358,0.00081987854,0.0000634738,0.0037627418],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99798906,0.00082836684,0.00015114964,0.0003781815,0.0005344382,0.000118832475],"domain_scores_gemma":[0.9932815,0.005338318,0.0003543064,0.00023253183,0.00066536886,0.00012807577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035263745,0.0006065582,0.0009414523,0.002537284,0.0009974835,0.0033419954,0.0018332874,0.0012868164,0.0037289292],"category_scores_gemma":[0.017617572,0.00055820745,0.0015758955,0.0016841352,0.0014237543,0.004692104,0.0011817304,0.0020096544,0.00061139767],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005799967,0.00022124933,0.0011979603,0.00079970417,0.00036397748,0.00031913465,0.001718189,0.4376766,0.0028261621,0.3030152,0.0064341743,0.24484763],"study_design_scores_gemma":[0.00003117016,0.000027399854,0.0002392945,0.000049041613,0.000040090355,0.00002340757,0.0001193329,0.8102223,0.0006199021,0.18739313,0.0012094268,0.000025445219],"about_ca_topic_score_codex":0.015009917,"about_ca_topic_score_gemma":0.013490825,"teacher_disagreement_score":0.015009917,"about_ca_system_score_codex":0.003330785,"about_ca_system_score_gemma":0.0017109644,"threshold_uncertainty_score":0.029845059},"labels":[],"label_agreement":null},{"id":"W17684","doi":"10.1111/j.2042-7158.1977.tb11305.x","title":"Word pairs in language modeling for information retrieval","year":2004,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Bigram; Vector space model; Computer science; Word (group theory); Language model; Natural language processing; Artificial intelligence; Constraint (computer-aided design); Term Discrimination; Adjacency list; Information retrieval; Visual Word; Algorithm; Mathematics","score_opus":0.02039087438294956,"score_gpt":0.2508621982835105,"score_spread":0.2304713239005609,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W17684","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007965537,0.0045934403,0.98074085,0.0013220921,0.00032223956,0.00013520622,0.00040768652,0.0009563682,0.0035566106],"genre_scores_gemma":[0.35417113,0.0070048682,0.61987627,0.0010125224,0.0011208863,0.0013421451,0.0023358134,0.0007862967,0.012350085],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99675983,0.002163726,0.00016923547,0.00041610407,0.00037690823,0.000114224684],"domain_scores_gemma":[0.9962185,0.002910208,0.00017940722,0.00030709835,0.00031286068,0.00007192469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033550907,0.0013390668,0.0018890612,0.0023013647,0.0010427465,0.002549868,0.0017222073,0.0017117744,0.007045361],"category_scores_gemma":[0.012543875,0.0006470784,0.001732095,0.003422169,0.00078124995,0.007051956,0.0017156687,0.0022077174,0.0036584663],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003056397,0.00022762983,0.0017073661,0.0008705329,0.00043841702,0.0003300948,0.0009555678,0.23070876,0.0035489046,0.4037711,0.018953426,0.33818245],"study_design_scores_gemma":[0.00003246954,0.000056373447,0.00024675517,0.000052665433,0.00008202162,0.00011461358,0.00012477498,0.66547936,0.0009818305,0.32121497,0.011562773,0.000051417624],"about_ca_topic_score_codex":0.005265391,"about_ca_topic_score_gemma":0.0043047825,"teacher_disagreement_score":0.007045361,"about_ca_system_score_codex":0.001271229,"about_ca_system_score_gemma":0.0013430101,"threshold_uncertainty_score":0.023569047},"labels":[],"label_agreement":null},{"id":"W1777006545","doi":"10.1111/j.1756-8765.2010.01108.x","title":"Comparing Methods for Single Paragraph Similarity Analysis","year":2010,"lang":"en","type":"article","venue":"Topics in Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada","funders":"Defence Research and Development Canada","keywords":"Paragraph; Computer science; Natural language processing; Artificial intelligence; Similarity (geometry); Word (group theory); Focus (optics); Semantic similarity; Vector space model; Simple (philosophy); Domain (mathematical analysis); Information retrieval; Linguistics; Mathematics; World Wide Web","score_opus":0.14923709677370298,"score_gpt":0.42624691451962937,"score_spread":0.27700981774592637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1777006545","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10798252,0.0054804534,0.8681693,0.00037181098,0.00046841693,0.00066649955,0.0021946295,0.008021568,0.0066447947],"genre_scores_gemma":[0.32710156,0.0012422265,0.6606227,0.0001415477,0.00026101482,0.00084651756,0.006040477,0.0016048729,0.0021390144],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.982064,0.00842347,0.0014637719,0.0030140989,0.004682335,0.00035235527],"domain_scores_gemma":[0.9249852,0.05467168,0.002117573,0.00910654,0.008285484,0.0008336186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013518387,0.0013691813,0.0014846921,0.014646399,0.0008740912,0.0034784754,0.0030906997,0.0021757828,0.0040604193],"category_scores_gemma":[0.07110843,0.0006224399,0.002519995,0.008431969,0.00072908466,0.0062677516,0.0029597767,0.0019526762,0.0028337608],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015329022,0.00046207494,0.021412596,0.0016034968,0.002235415,0.0001418147,0.0011209073,0.03660296,0.0075680763,0.009278634,0.0125648845,0.9054762],"study_design_scores_gemma":[0.00028688886,0.000898412,0.02317914,0.0002968027,0.0006230826,0.00079576334,0.0017115184,0.89411324,0.018123282,0.044066455,0.015635798,0.00026954518],"about_ca_topic_score_codex":0.0022316205,"about_ca_topic_score_gemma":0.0032504392,"teacher_disagreement_score":0.014646399,"about_ca_system_score_codex":0.0012535517,"about_ca_system_score_gemma":0.0011978585,"threshold_uncertainty_score":0.07149297},"labels":[],"label_agreement":null},{"id":"W180888214","doi":"","title":"Evaluating Distributional Models of Semantics for Syntactically Invariant Inference","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Distributional semantics; Phrase; Inference; Natural language processing; Invariant (physics); Artificial intelligence; Lemma (botany); Semantics (computer science); Sentence; Focus (optics); Mathematics; Programming language","score_opus":0.17678125876358503,"score_gpt":0.3826369719767289,"score_spread":0.2058557132131439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W180888214","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17629302,0.0006915925,0.814588,0.00112756,0.00008615578,0.00014473194,0.0005041534,0.002180626,0.0043840813],"genre_scores_gemma":[0.8559792,0.00026236608,0.13997476,0.0001918823,0.00011203653,0.00014606115,0.0019417192,0.00029745826,0.0010945617],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9917203,0.005185267,0.00043466853,0.0011399008,0.001209975,0.0003098601],"domain_scores_gemma":[0.9634428,0.029227303,0.0010387191,0.0035629265,0.0020436647,0.000684542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014204153,0.001104513,0.0015169465,0.0023342064,0.00096657564,0.0037252496,0.0022823752,0.001894108,0.0033425223],"category_scores_gemma":[0.046668917,0.0005550914,0.0015913465,0.0017378671,0.0017926636,0.011265712,0.0033769193,0.0031979096,0.0009221156],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018542685,0.00083921873,0.022206524,0.0007667299,0.00065297115,0.000260198,0.0010487034,0.36831003,0.010909364,0.14792974,0.006236262,0.438986],"study_design_scores_gemma":[0.000022141825,0.00008942528,0.0007501675,0.00001864077,0.00003380578,0.000039812814,0.000083220955,0.9350096,0.0019377441,0.061617456,0.00038077307,0.000017085858],"about_ca_topic_score_codex":0.003982431,"about_ca_topic_score_gemma":0.0064362446,"teacher_disagreement_score":0.014204153,"about_ca_system_score_codex":0.0021784035,"about_ca_system_score_gemma":0.0017264172,"threshold_uncertainty_score":0.075119674},"labels":[],"label_agreement":null},{"id":"W1814922127","doi":"","title":"Summarizing Blog Entries versus News Texts","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Automatic summarization; Blogosphere; Computer science; Information retrieval; Popularity; Social media; Categorization; World Wide Web; Variety (cybernetics); Event (particle physics); Multi-document summarization; Data science; Artificial intelligence; The Internet; Psychology","score_opus":0.029729404170158878,"score_gpt":0.261370563285352,"score_spread":0.23164115911519315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1814922127","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79996145,0.0028727315,0.17481883,0.0007394655,0.0004638043,0.0005098141,0.0076520815,0.0065349154,0.0064469473],"genre_scores_gemma":[0.83725196,0.0014129299,0.13962881,0.00013967279,0.00036363635,0.00021615208,0.016186094,0.0005997729,0.0042010546],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99856097,0.00043101123,0.00022230143,0.00032818326,0.0003935751,0.00006408384],"domain_scores_gemma":[0.9902208,0.0054015354,0.0011054634,0.00084487343,0.0022677132,0.00015955773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017067996,0.0006646801,0.0006922629,0.002118905,0.00037664524,0.0014563404,0.00043083672,0.00045306844,0.0013992245],"category_scores_gemma":[0.017400464,0.00020289466,0.00030837068,0.002337033,0.00017998903,0.0016162174,0.0005903502,0.0003859017,0.0006493095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028156962,0.00038385385,0.042801354,0.0027612105,0.00048407287,0.0006097404,0.005306957,0.02721886,0.09819998,0.0021227445,0.022916913,0.7943787],"study_design_scores_gemma":[0.0004263741,0.004994721,0.14524679,0.0005490674,0.0025153712,0.0013758814,0.011066964,0.48567525,0.25176847,0.014215568,0.08180185,0.0003638536],"about_ca_topic_score_codex":0.0015507916,"about_ca_topic_score_gemma":0.0029931972,"teacher_disagreement_score":0.002118905,"about_ca_system_score_codex":0.00021891521,"about_ca_system_score_gemma":0.00029220994,"threshold_uncertainty_score":0.009026468},"labels":[],"label_agreement":null},{"id":"W1815301076","doi":"","title":"Measuring the Non-compositionality of Multiword Expressions","year":2010,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Principle of compositionality; Computer science; Artificial intelligence; Natural language processing; Semantics (computer science); Metric (unit); Natural language; Question answering; Combinatory categorial grammar; Distributional semantics; Expression (computer science); Information extraction; Programming language; Semantic similarity; Link grammar; Rule-based machine translation","score_opus":0.0730969691908334,"score_gpt":0.32514661388054306,"score_spread":0.25204964468970964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1815301076","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4840883,0.0015936336,0.5077538,0.00024741516,0.00015708826,0.00017966494,0.0005141906,0.0011574134,0.0043085963],"genre_scores_gemma":[0.8126592,0.0005931136,0.18261057,0.00008799208,0.00009338513,0.00016966919,0.0018266233,0.0003244131,0.001634966],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962781,0.001206149,0.0005335329,0.00095445727,0.00087545,0.00015233975],"domain_scores_gemma":[0.98637044,0.008462582,0.0015140852,0.0014831077,0.0018333149,0.00033644974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024562236,0.00074496906,0.0007226532,0.0028273195,0.0008347637,0.001822397,0.0008609626,0.0011275944,0.0013259294],"category_scores_gemma":[0.020456329,0.00041790915,0.00062590925,0.0019135417,0.00088648585,0.0048124897,0.0026679682,0.0010338421,0.0010175563],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00095899747,0.000299506,0.052718654,0.001413574,0.00039297243,0.0006661928,0.0022997411,0.012379215,0.29919678,0.016817246,0.001576804,0.61128026],"study_design_scores_gemma":[0.0000940984,0.0012326004,0.10833195,0.00023426262,0.00049042504,0.0036937569,0.0033801463,0.5285692,0.24974047,0.077282645,0.02670508,0.00024537242],"about_ca_topic_score_codex":0.00048230393,"about_ca_topic_score_gemma":0.00083859364,"teacher_disagreement_score":0.0028273195,"about_ca_system_score_codex":0.00038023363,"about_ca_system_score_gemma":0.00059572665,"threshold_uncertainty_score":0.012989879},"labels":[],"label_agreement":null},{"id":"W1818810666","doi":"","title":"Positional Language Models for Clinical Information Retrieval","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"","keywords":"Computer science; Information retrieval; Weighting; Language model; Data collection; Representation (politics); Natural language processing; Test (biology); Data mining; Artificial intelligence; Statistics; Mathematics","score_opus":0.04293954222260586,"score_gpt":0.341099770726688,"score_spread":0.29816022850408214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1818810666","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014222171,0.003148269,0.96877944,0.0019212385,0.00014639001,0.00044620788,0.0045377985,0.0025904675,0.0042079706],"genre_scores_gemma":[0.45951664,0.0046881074,0.5079332,0.0010875739,0.00066887954,0.0026283294,0.012899082,0.00054812065,0.010030051],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964552,0.0020782943,0.00028864635,0.00048142928,0.0005455337,0.00015089201],"domain_scores_gemma":[0.98965454,0.008398191,0.0005305499,0.00053207274,0.00075181987,0.00013277664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004533331,0.001135044,0.0010789082,0.005511851,0.0008468925,0.0029073115,0.002085603,0.0016126197,0.0059993123],"category_scores_gemma":[0.018983454,0.00071400625,0.0019560661,0.0051840045,0.0007116734,0.0042272015,0.0012276828,0.0014937132,0.0034639898],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006478382,0.00036261996,0.005013021,0.0014900108,0.0006727351,0.0004551498,0.0018826413,0.26949063,0.0046685003,0.23533867,0.031566985,0.44841123],"study_design_scores_gemma":[0.000072089504,0.00008135978,0.00084530946,0.000073868905,0.000118845644,0.00015169756,0.00016049495,0.85012877,0.00093952485,0.13587043,0.011507188,0.000050444596],"about_ca_topic_score_codex":0.017235337,"about_ca_topic_score_gemma":0.01820262,"teacher_disagreement_score":0.017235337,"about_ca_system_score_codex":0.002427772,"about_ca_system_score_gemma":0.002572841,"threshold_uncertainty_score":0.03427005},"labels":[],"label_agreement":null},{"id":"W1828724394","doi":"10.48550/arxiv.1410.2455","title":"BilBOWA: Fast Bilingual Distributed Representations without Word Alignments","year":2014,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":314,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Canadian Institute for Advanced Research","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Task (project management); Word (group theory); Bilingual dictionary; Sentence; Feature (linguistics); Set (abstract data type); Machine translation; Translation (biology); Speech recognition; Linguistics","score_opus":0.055720460782726695,"score_gpt":0.20385800545450294,"score_spread":0.14813754467177626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1828724394","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017020863,0.0005784828,0.9694914,0.00026668064,0.00018479866,0.00010258713,0.0014693366,0.009184909,0.0017008571],"genre_scores_gemma":[0.26966375,0.00071097846,0.6981762,0.0005450118,0.0002919687,0.00089359115,0.01769568,0.0019558019,0.010067044],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991549,0.0002739545,0.000046921257,0.00027944648,0.0001474419,0.000097322],"domain_scores_gemma":[0.9991197,0.00023644434,0.00007184264,0.00034500455,0.00015071285,0.000076416894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014231482,0.0015790955,0.001329589,0.0014143538,0.0006224874,0.0014892543,0.002293299,0.0011059084,0.0061391336],"category_scores_gemma":[0.0043869205,0.00067928364,0.0011588803,0.0019022351,0.0005303894,0.0039947475,0.0033998839,0.0023131203,0.007450167],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058435905,0.00044734622,0.0025671164,0.0002916734,0.00024272829,0.00016068609,0.0002838374,0.08080888,0.014446697,0.023496488,0.046213206,0.8304571],"study_design_scores_gemma":[0.000086708715,0.00013392708,0.0005172955,0.000031891876,0.000038040933,0.00012931839,0.00008422535,0.9384073,0.006106061,0.041559216,0.012867432,0.000038619288],"about_ca_topic_score_codex":0.0035715085,"about_ca_topic_score_gemma":0.0066402797,"teacher_disagreement_score":0.0061391336,"about_ca_system_score_codex":0.00061799237,"about_ca_system_score_gemma":0.0017582236,"threshold_uncertainty_score":0.020537436},"labels":[],"label_agreement":null},{"id":"W183308553","doi":"","title":"Evaluating WordNet Features in Text Classification Models.","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina","funders":"","keywords":"WordNet; Computer science; Artificial intelligence; Naive Bayes classifier; Classifier (UML); Support vector machine; Natural language processing; Machine learning; Information retrieval","score_opus":0.09596409055603045,"score_gpt":0.3292957080679671,"score_spread":0.23333161751193665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W183308553","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61937845,0.009168321,0.3333677,0.0028792508,0.0012781607,0.0018052359,0.0071761836,0.008951659,0.015995122],"genre_scores_gemma":[0.8489249,0.0011894424,0.13437186,0.00032692152,0.00029422418,0.0006967266,0.010483626,0.00020398914,0.003508408],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99433494,0.002515679,0.0005530552,0.00081988587,0.0015289915,0.00024748078],"domain_scores_gemma":[0.973532,0.021420093,0.0012126487,0.00097475393,0.0024937973,0.00036668705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010571971,0.002627467,0.0014213211,0.007272211,0.0008233151,0.0028360945,0.0012841038,0.0023004517,0.0023251923],"category_scores_gemma":[0.029375454,0.00041021354,0.0012082729,0.0043879263,0.00057229726,0.006767552,0.0012093673,0.0015867385,0.0015634961],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015262767,0.0017497614,0.057943765,0.0011616807,0.0012526693,0.0003759316,0.00042100257,0.26625857,0.00619187,0.005198351,0.02088488,0.6370353],"study_design_scores_gemma":[0.00006208779,0.00035801672,0.0032944889,0.00006813146,0.000116843716,0.00007369143,0.00015694217,0.98614746,0.003191901,0.0043732114,0.0021344982,0.000022764296],"about_ca_topic_score_codex":0.007129029,"about_ca_topic_score_gemma":0.011516544,"teacher_disagreement_score":0.010571971,"about_ca_system_score_codex":0.0020690793,"about_ca_system_score_gemma":0.0013676469,"threshold_uncertainty_score":0.055910587},"labels":[],"label_agreement":null},{"id":"W1838551284","doi":"10.1136/amiajnl-2013-002214","title":"Natural language processing: algorithms and tools to extract computable information from EHRs and from the biomedical literature","year":2013,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"CONTEST; Computer science; Narrative; Natural language processing; Health records; Artificial intelligence; Task (project management); Interpretation (philosophy); Structuring; Information retrieval; Biomedical text mining; Quality (philosophy); Text mining; Linguistics; Programming language; Health care","score_opus":0.005966570501120155,"score_gpt":0.2444841220209323,"score_spread":0.23851755151981216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1838551284","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002969557,0.007898379,0.97111815,0.0026624014,0.0003081642,0.0006808842,0.0055123563,0.005968153,0.0028819975],"genre_scores_gemma":[0.014698185,0.006084793,0.9666205,0.00049179746,0.0003893248,0.0010231477,0.009229584,0.00033685155,0.0011258849],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99315417,0.0028414505,0.001171477,0.001002677,0.0017096766,0.00012045676],"domain_scores_gemma":[0.9807818,0.015633143,0.00095763017,0.0011254764,0.001337596,0.00016440953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009997256,0.0025283862,0.001969578,0.015740028,0.0013631015,0.0066200793,0.0029554204,0.0019230363,0.004546032],"category_scores_gemma":[0.02826357,0.000962064,0.003168588,0.012426809,0.0020284236,0.0081508495,0.0036797586,0.0030814428,0.004501986],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011605349,0.00022154932,0.0020939503,0.004864482,0.00038767632,0.00039502897,0.001068562,0.009240581,0.0053385445,0.04299458,0.04104234,0.89223665],"study_design_scores_gemma":[0.00013541988,0.0001813278,0.0042957803,0.0026654275,0.00040119144,0.0011608998,0.0016951403,0.24334776,0.013674315,0.51276857,0.21938379,0.0002903926],"about_ca_topic_score_codex":0.0039614085,"about_ca_topic_score_gemma":0.0045135194,"teacher_disagreement_score":0.015740028,"about_ca_system_score_codex":0.0016819183,"about_ca_system_score_gemma":0.003844746,"threshold_uncertainty_score":0.052871227},"labels":[],"label_agreement":null},{"id":"W1846406155","doi":"10.19173/irrodl.v13i5.1269","title":"Instructor-aided asynchronous question answering system for online education and distance learning","year":2012,"lang":"en","type":"article","venue":"The International Review of Research in Open and Distributed Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Distance education; Asynchronous communication; Computer science; Question answering; Exploit; Synchronous learning; Perspective (graphical); Learning Management; Asynchronous learning; Online learning; Educational technology; Quality (philosophy); Artificial intelligence; Multimedia; World Wide Web; Mathematics education; Teaching method; Cooperative learning; Psychology","score_opus":0.06494774927665814,"score_gpt":0.42353235705436787,"score_spread":0.35858460777770973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1846406155","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09643244,0.00046900698,0.7700203,0.00069901795,0.00025849338,0.0014315041,0.002865293,0.114034064,0.013789853],"genre_scores_gemma":[0.41999218,0.00023393851,0.53935486,0.00045607513,0.00032057753,0.0012771958,0.0077053197,0.0011261017,0.029533762],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989335,0.00035189086,0.00009422078,0.0003560363,0.00020209243,0.00006233603],"domain_scores_gemma":[0.9959268,0.0017842762,0.00023133474,0.00071211957,0.00085174467,0.0004936605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002377928,0.00058517285,0.0006914677,0.0011195196,0.00066215004,0.001125905,0.0016216147,0.0013417873,0.020328546],"category_scores_gemma":[0.005389146,0.00024792048,0.00040297568,0.0007340729,0.00020504344,0.0021174124,0.0013078419,0.0008798108,0.010073045],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014967648,0.0015059138,0.0065499754,0.00045679059,0.00006983045,0.0004631515,0.0013938881,0.002254528,0.06571928,0.0048275813,0.048666894,0.8665954],"study_design_scores_gemma":[0.0011936194,0.0032393907,0.021851577,0.00015536134,0.0003704397,0.0016316618,0.0009369326,0.3805979,0.13901974,0.013500554,0.43723586,0.0002669535],"about_ca_topic_score_codex":0.0008570391,"about_ca_topic_score_gemma":0.00092425395,"teacher_disagreement_score":0.020328546,"about_ca_system_score_codex":0.00042730098,"about_ca_system_score_gemma":0.00076430326,"threshold_uncertainty_score":0.0680058},"labels":[],"label_agreement":null},{"id":"W1849027473","doi":"10.1007/978-3-642-38824-8_11","title":"Evaluating Syntactic Sentence Compression for Text Summarisation","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Natural language processing; Sentence; Computer science; Linguistics; Artificial intelligence; Philosophy","score_opus":0.06215909734522565,"score_gpt":0.316962817238673,"score_spread":0.2548037198934473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1849027473","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40183306,0.017975371,0.5120078,0.0014294975,0.0022953914,0.001335629,0.008777536,0.03504749,0.019298222],"genre_scores_gemma":[0.51965034,0.003922165,0.419338,0.00034185924,0.0015598417,0.0007273368,0.039891254,0.0015584959,0.013010736],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983164,0.00057862175,0.00016533755,0.0002773947,0.00052432436,0.00013796183],"domain_scores_gemma":[0.99388146,0.004047351,0.00025303132,0.00033662305,0.0013148938,0.00016662135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020191013,0.0022476204,0.0017794814,0.002354843,0.0007321778,0.0017766672,0.0012297033,0.0016271485,0.009049553],"category_scores_gemma":[0.0077321352,0.00047000268,0.0008477612,0.0021459952,0.00035768244,0.0021856138,0.0012420346,0.0012589098,0.005138219],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016867431,0.0003385745,0.0012614811,0.0012039703,0.0002939054,0.0003400731,0.0002834137,0.016354613,0.056087986,0.0011292733,0.029283492,0.8917364],"study_design_scores_gemma":[0.0006133775,0.0027193178,0.008794719,0.00023875956,0.0010338419,0.0008633283,0.00095000764,0.8317289,0.12474718,0.008062844,0.020102784,0.00014500848],"about_ca_topic_score_codex":0.0025022095,"about_ca_topic_score_gemma":0.0036050277,"teacher_disagreement_score":0.009049553,"about_ca_system_score_codex":0.0005452465,"about_ca_system_score_gemma":0.0009515727,"threshold_uncertainty_score":0.030273795},"labels":[],"label_agreement":null},{"id":"W1871378960","doi":"","title":"Extending the Entity-based Coherence Model with Multiple Ranks","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Coreference; Coherence (philosophical gambling strategy); Computer science; Pairwise comparison; Artificial intelligence; Natural language processing; Sentence; Training set; Set (abstract data type); Component (thermodynamics); Information retrieval; Statistics; Resolution (logic); Mathematics; Programming language","score_opus":0.0360831324396391,"score_gpt":0.24706129000985785,"score_spread":0.21097815757021876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1871378960","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057263605,0.000733612,0.93434536,0.0008261636,0.00007481408,0.00015887049,0.00093063107,0.0019011588,0.0037658059],"genre_scores_gemma":[0.78322065,0.00058372354,0.20459907,0.00041678187,0.00026164204,0.00026950962,0.0032962724,0.00041112676,0.006941153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99789774,0.0008571067,0.00014290324,0.00052950624,0.00041827516,0.00015453144],"domain_scores_gemma":[0.9947666,0.002912567,0.0004289335,0.0008808995,0.0008546492,0.00015632986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037097686,0.00089758524,0.0008736211,0.0016287257,0.00055613654,0.0015925263,0.002084702,0.0012869402,0.0039330265],"category_scores_gemma":[0.012976591,0.00061474694,0.0010934595,0.002211139,0.0006327215,0.005799188,0.0017122128,0.0017953031,0.0014947754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008115942,0.0004674227,0.019056302,0.00044752832,0.00064101163,0.00044709703,0.0012736343,0.3910449,0.010830296,0.043738343,0.012899485,0.5183423],"study_design_scores_gemma":[0.00006454068,0.00020621078,0.002868497,0.000023398083,0.000121302204,0.0001347943,0.000054456414,0.97301346,0.0016130896,0.017410705,0.0044350997,0.000054425887],"about_ca_topic_score_codex":0.007857859,"about_ca_topic_score_gemma":0.013468004,"teacher_disagreement_score":0.007857859,"about_ca_system_score_codex":0.0006491548,"about_ca_system_score_gemma":0.001087819,"threshold_uncertainty_score":0.019619346},"labels":[],"label_agreement":null},{"id":"W1879966306","doi":"","title":"Long Short-Term Memory Over Recursive Structures","year":2015,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":273,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Ottawa; National Research Council Canada","funders":"","keywords":"Computer science; Leverage (statistics); Artificial intelligence; Parsing; Theoretical computer science; Natural language processing","score_opus":0.03825664726599463,"score_gpt":0.27864304625467984,"score_spread":0.24038639898868522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1879966306","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07111851,0.0024780177,0.91352266,0.0006922918,0.00014442498,0.000049544447,0.0006518159,0.0032657548,0.0080770645],"genre_scores_gemma":[0.8760471,0.0019984653,0.11362667,0.00020879389,0.00012360558,0.00010828264,0.0010267977,0.00017912289,0.006681083],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998429,0.000038378963,0.000012013036,0.000062595354,0.000021848353,0.000022298105],"domain_scores_gemma":[0.9995536,0.00024279827,0.000052593703,0.000066459885,0.000064139036,0.000020387573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00034259167,0.00072020944,0.00043902692,0.0003954242,0.00027794007,0.00068509363,0.0011005695,0.0008038584,0.0036859612],"category_scores_gemma":[0.0018194937,0.00029277906,0.0004237522,0.0007259872,0.000537596,0.003171249,0.00067058345,0.0010557581,0.0009293194],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039087576,0.00008687101,0.001696897,0.000607393,0.00013102478,0.00041394462,0.00042154416,0.37293637,0.043803915,0.12945986,0.009376834,0.44067442],"study_design_scores_gemma":[0.000010864805,0.000057503494,0.00035218324,0.000028485847,0.000031976288,0.0000580734,0.000025558547,0.9413443,0.0059736096,0.04972975,0.0023758002,0.000011958926],"about_ca_topic_score_codex":0.0046101604,"about_ca_topic_score_gemma":0.006161918,"teacher_disagreement_score":0.0046101604,"about_ca_system_score_codex":0.0007314786,"about_ca_system_score_gemma":0.00062667596,"threshold_uncertainty_score":0.012330711},"labels":[],"label_agreement":null},{"id":"W1881423045","doi":"10.1007/978-3-642-13059-5_54","title":"Peer-Based Intelligent Tutoring Systems: A Corpus-Oriented Approach","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Multimedia; World Wide Web; Human–computer interaction","score_opus":0.029749265761990645,"score_gpt":0.24785607824101968,"score_spread":0.21810681247902905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1881423045","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036161263,0.0034039025,0.9053404,0.0019300343,0.00037423545,0.0009973123,0.002760126,0.006213735,0.042819016],"genre_scores_gemma":[0.45444465,0.0030620391,0.50199705,0.00033666083,0.0007339998,0.0017328651,0.0072576893,0.0019833727,0.028451623],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960836,0.0021638733,0.00019994804,0.000520193,0.00091940665,0.000112940266],"domain_scores_gemma":[0.98900944,0.007481955,0.00025071678,0.0010135135,0.0019318827,0.00031246128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00472696,0.00092498685,0.0017922992,0.003784066,0.0022856721,0.004824449,0.004564911,0.0019966073,0.010894489],"category_scores_gemma":[0.02173476,0.0007346723,0.00058144814,0.0049430556,0.001669998,0.008740215,0.0038172544,0.001647794,0.0036514562],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005181064,0.00052424415,0.0037797613,0.0021843247,0.00026438726,0.00067555124,0.0070196157,0.066826865,0.015698005,0.10801908,0.045480616,0.74900943],"study_design_scores_gemma":[0.00018703248,0.00040382182,0.0033366873,0.00034989687,0.00041697937,0.00096949993,0.004654319,0.68179333,0.023380818,0.13361776,0.15070409,0.00018577033],"about_ca_topic_score_codex":0.0034796703,"about_ca_topic_score_gemma":0.005687301,"teacher_disagreement_score":0.010894489,"about_ca_system_score_codex":0.0015436743,"about_ca_system_score_gemma":0.0021170196,"threshold_uncertainty_score":0.036445677},"labels":[],"label_agreement":null},{"id":"W1883367367","doi":"10.1007/978-3-540-87391-4_1","title":"The Future of Text-Meaning in Computational Linguistics","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Meaning (existential); Computer science; Perspective (graphical); Computational linguistics; Linguistics; Artificial intelligence; Natural language processing; Epistemology; Philosophy","score_opus":0.017460770707665493,"score_gpt":0.24569104801959635,"score_spread":0.22823027731193085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1883367367","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014928627,0.36394608,0.38759932,0.08970595,0.007212739,0.00006675447,0.00073882675,0.0010505263,0.13475117],"genre_scores_gemma":[0.57050925,0.12737165,0.22274438,0.007462631,0.021440445,0.00029974798,0.0013144624,0.0013507055,0.047506787],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9978855,0.0013080578,0.0001353927,0.0002752261,0.00030866516,0.000087088745],"domain_scores_gemma":[0.98955345,0.009007348,0.00020124241,0.0006304294,0.00041276138,0.00019477126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039643953,0.0007581515,0.0012131967,0.0033051923,0.0015722188,0.00897675,0.0016528127,0.0026142125,0.009298409],"category_scores_gemma":[0.009912397,0.0007421358,0.0007881165,0.0049217516,0.01117243,0.033623457,0.002385038,0.0043253493,0.0021408063],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017926808,0.0000113442275,0.00011857822,0.00032353812,0.000010543139,0.00003302132,0.0011822135,0.0003524861,0.0002124909,0.93951994,0.008949133,0.04926885],"study_design_scores_gemma":[0.000005383582,0.0000048528887,0.00009173432,0.000103962884,0.0000048792895,0.000048553356,0.00022922442,0.0015394758,0.00009945652,0.95936006,0.038503636,0.000008830265],"about_ca_topic_score_codex":0.0011644723,"about_ca_topic_score_gemma":0.0013789369,"teacher_disagreement_score":0.009298409,"about_ca_system_score_codex":0.00207502,"about_ca_system_score_gemma":0.0014585693,"threshold_uncertainty_score":0.031106293},"labels":[],"label_agreement":null},{"id":"W1902272523","doi":"10.1002/chp.21190","title":"A Community of Practice for Knowledge Translation Trainees: An Innovative Approach for Learning and Collaboration","year":2013,"lang":"en","type":"article","venue":"Journal of Continuing Education in the Health Professions","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Kelowna General Hospital; St. Michael's Hospital; Institute for Clinical Evaluative Sciences; McGill University; Capital District Health Authority; Ottawa Hospital; Interior Health; Hospital for Sick Children; University of British Columbia; SickKids Foundation; University of British Columbia, Okanagan Campus; Centre hospitalier universitaire de Québec; Nova Scotia Health Authority; Dalhousie University","funders":"","keywords":"Knowledge translation; Knowledge management; Community of practice; Knowledge sharing; Medical education; Professional development; Field (mathematics); Diversity (politics); Continuing professional development; Computer science; Medicine; Psychology; Sociology; Pedagogy","score_opus":0.10415007078679604,"score_gpt":0.454720237275189,"score_spread":0.35057016648839295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1902272523","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10624767,0.0007997546,0.72046804,0.038043085,0.0007769271,0.0064893803,0.00014313601,0.0012778932,0.12575412],"genre_scores_gemma":[0.43771452,0.000363784,0.53768545,0.0015670774,0.00016768994,0.0047182757,0.00011457118,0.00018921652,0.017479414],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.937676,0.05207917,0.0011029566,0.0038761608,0.0034390744,0.0018266442],"domain_scores_gemma":[0.95099175,0.023914807,0.002573631,0.00696544,0.0044235196,0.011130854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03597395,0.0009557696,0.0007480027,0.0045263236,0.017901912,0.017107576,0.006267838,0.006533512,0.012596996],"category_scores_gemma":[0.041891653,0.0009787048,0.0016382134,0.0026926256,0.015252589,0.019303247,0.033174943,0.004517107,0.0025365672],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024287667,0.001594528,0.0071280934,0.0008125254,0.00007886303,0.0011987237,0.32662067,0.0019365563,0.002396429,0.3974259,0.015451997,0.24511291],"study_design_scores_gemma":[0.00047254216,0.0013675235,0.0025848232,0.0014990615,0.0001125195,0.0019973433,0.22853932,0.020865258,0.001697783,0.35995102,0.38067603,0.00023680402],"about_ca_topic_score_codex":0.002516816,"about_ca_topic_score_gemma":0.004853742,"teacher_disagreement_score":0.03597395,"about_ca_system_score_codex":0.007555954,"about_ca_system_score_gemma":0.024571864,"threshold_uncertainty_score":0.1902507},"labels":[],"label_agreement":null},{"id":"W191017350","doi":"","title":"Enlarging Paraphrase Collections through Generalization and Instantiation","year":2012,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Japan Society for the Promotion of Science; National Research Council Canada","keywords":"Paraphrase; Computer science; Natural language processing; Generalization; Artificial intelligence; Quality (philosophy); Parallel corpora; Exploit; Machine translation; Mathematics","score_opus":0.02717435825852078,"score_gpt":0.255400939592151,"score_spread":0.22822658133363022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W191017350","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.090509154,0.0013766694,0.881353,0.0005413157,0.000056581513,0.0011209024,0.0016110368,0.0146516785,0.00877957],"genre_scores_gemma":[0.2611343,0.00083585625,0.7262029,0.00023157979,0.00009580914,0.0004939243,0.0056246156,0.001206482,0.0041744756],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99703515,0.00093526515,0.0003383172,0.00094495196,0.0005878348,0.0001585235],"domain_scores_gemma":[0.9903349,0.0043320674,0.00047578837,0.0038451247,0.00079476583,0.00021740177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033716322,0.0014816901,0.0016212194,0.006173621,0.0010152967,0.0019158764,0.0026470353,0.0011230547,0.0050166966],"category_scores_gemma":[0.012725355,0.0010877324,0.0016709603,0.004872291,0.0010052382,0.0059865126,0.0043252283,0.0020195039,0.0028156783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025852278,0.0005588956,0.006496619,0.0007464669,0.0002548295,0.0008150338,0.0018443898,0.009824048,0.05569328,0.009741295,0.012242355,0.9015243],"study_design_scores_gemma":[0.00019330466,0.00077741256,0.022093441,0.00047063053,0.0008270806,0.0051827007,0.0024757525,0.65764976,0.123640835,0.08193064,0.10448211,0.00027631197],"about_ca_topic_score_codex":0.0013715076,"about_ca_topic_score_gemma":0.003431099,"teacher_disagreement_score":0.006173621,"about_ca_system_score_codex":0.00062445394,"about_ca_system_score_gemma":0.00095164444,"threshold_uncertainty_score":0.017831087},"labels":[],"label_agreement":null},{"id":"W1921587136","doi":"10.1007/s10579-015-9318-3","title":"Cross level semantic similarity: an evaluation framework for universal measures of similarity","year":2015,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"European Research Council","keywords":"Computer science; Natural language processing; Similarity (geometry); Semantic similarity; Sentence; Task (project management); Artificial intelligence; Word (group theory); SemEval; Paragraph; Process (computing); Meaning (existential); WordNet; Information retrieval; Linguistics; Psychology; World Wide Web","score_opus":0.21232820109561595,"score_gpt":0.3869887300338786,"score_spread":0.17466052893826267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1921587136","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021688528,0.0010031083,0.97204685,0.00014308491,0.000088809225,0.0004102549,0.0006332746,0.0013459966,0.0026400108],"genre_scores_gemma":[0.40879327,0.00055454083,0.58532137,0.00014636011,0.00019346453,0.0011357986,0.0019179165,0.0005842705,0.0013531148],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97493875,0.01090835,0.0027606743,0.002648959,0.008026661,0.0007165587],"domain_scores_gemma":[0.9645081,0.019548161,0.0023981351,0.0063810423,0.0059709023,0.0011935553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024129214,0.0013911831,0.0020710602,0.012123618,0.0014304467,0.0053690686,0.0024734442,0.0021485167,0.0029954636],"category_scores_gemma":[0.06528036,0.00052313606,0.001974848,0.007821052,0.0019568354,0.011363664,0.005723787,0.002126443,0.00087291974],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019071632,0.00091230817,0.028248578,0.0016784037,0.0017763411,0.00023768007,0.0026437575,0.028456664,0.019701356,0.18722643,0.010013435,0.7171978],"study_design_scores_gemma":[0.00020196753,0.0020268126,0.022483358,0.00056112825,0.001246642,0.0011793971,0.0015848479,0.5581739,0.029200079,0.3625684,0.020391693,0.00038172404],"about_ca_topic_score_codex":0.001939076,"about_ca_topic_score_gemma":0.0018627377,"teacher_disagreement_score":0.024129214,"about_ca_system_score_codex":0.0019513459,"about_ca_system_score_gemma":0.0019573811,"threshold_uncertainty_score":0.12760901},"labels":[],"label_agreement":null},{"id":"W192510948","doi":"10.1145/2567948.2579705","title":"Entity linking with a unified semantic representation","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Referent; Entity linking; Semantics (computer science); Representation (politics); Information retrieval; Graph; Focus (optics); Natural language processing; Knowledge base; Artificial intelligence; Theoretical computer science; Linguistics; Programming language","score_opus":0.024851335809388315,"score_gpt":0.25298859766457005,"score_spread":0.22813726185518174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W192510948","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004598072,0.0008613918,0.97792006,0.0004402982,0.00010898462,0.00015035876,0.0020106137,0.009436174,0.0044740485],"genre_scores_gemma":[0.16235055,0.0020785006,0.80769855,0.0005065066,0.00035079097,0.00039771397,0.016296227,0.0011753546,0.009145695],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974655,0.0006354773,0.00021785607,0.00090931915,0.0006274652,0.0001443554],"domain_scores_gemma":[0.9967096,0.0011625913,0.0003198666,0.0011733331,0.0005454404,0.00008914929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024911263,0.0013775969,0.0014597846,0.010224529,0.0011909384,0.00443354,0.0029472779,0.0026163096,0.0065007545],"category_scores_gemma":[0.009813102,0.0008223699,0.002408348,0.013932455,0.0008044611,0.011709244,0.0035892052,0.0026583613,0.005260064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019099111,0.00030896638,0.002206414,0.0007529143,0.00034151907,0.00059635646,0.00090796803,0.078084946,0.006214101,0.10853065,0.053609114,0.74825597],"study_design_scores_gemma":[0.000036921825,0.00006777841,0.0011128939,0.00022125902,0.00036541247,0.0007012317,0.00037718995,0.72448593,0.008672068,0.15988593,0.10396128,0.000112104026],"about_ca_topic_score_codex":0.004508577,"about_ca_topic_score_gemma":0.006361581,"teacher_disagreement_score":0.010224529,"about_ca_system_score_codex":0.0012886389,"about_ca_system_score_gemma":0.0021618786,"threshold_uncertainty_score":0.021747172},"labels":[],"label_agreement":null},{"id":"W1940242918","doi":"10.1145/2740908.2745403","title":"Question Classification by Approximating Semantics","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Semantics (computer science); Computer science; Task (project management); Semantic similarity; Natural language processing; Artificial intelligence; Distributional semantics; Theoretical computer science; Information retrieval; Programming language","score_opus":0.06951154520707058,"score_gpt":0.28784721707262384,"score_spread":0.21833567186555325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1940242918","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08444818,0.0010421723,0.90846306,0.0014414157,0.00006682442,0.00010750335,0.00032839508,0.0015500109,0.0025525314],"genre_scores_gemma":[0.748859,0.00044912292,0.24698001,0.00036425705,0.00016309369,0.00021145945,0.0016064381,0.00017151017,0.0011951544],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99447024,0.0026403596,0.00039829023,0.0013109922,0.0009049185,0.00027520303],"domain_scores_gemma":[0.9767662,0.016909773,0.001234386,0.0033018761,0.0015330102,0.00025470514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057871044,0.00087704987,0.0013801319,0.005280918,0.001059535,0.0038030886,0.0021838052,0.00208304,0.0022886698],"category_scores_gemma":[0.04929926,0.0005150312,0.0014497677,0.003838348,0.00241788,0.009452461,0.0033447742,0.0020865132,0.0010140085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060647615,0.00027324725,0.017227108,0.00061472185,0.00020994675,0.00014966632,0.002889514,0.1593702,0.006882871,0.22907417,0.009018535,0.5736835],"study_design_scores_gemma":[0.00002696469,0.000048300302,0.0013627209,0.000044037857,0.000039320483,0.000091955735,0.00030632055,0.6877296,0.0022566335,0.30469072,0.0033802325,0.000023081064],"about_ca_topic_score_codex":0.0038482575,"about_ca_topic_score_gemma":0.0024001822,"teacher_disagreement_score":0.0057871044,"about_ca_system_score_codex":0.0020943845,"about_ca_system_score_gemma":0.0013682819,"threshold_uncertainty_score":0.030605495},"labels":[],"label_agreement":null},{"id":"W196214544","doi":"","title":"Generating Text with Recurrent Neural Networks","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1174,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of New Brunswick","funders":"","keywords":"Recurrent neural network; Computer science; Language model; Sequence (biology); Artificial intelligence; Hessian matrix; Artificial neural network; Multiplicative function; Machine learning","score_opus":0.04898282774934054,"score_gpt":0.2321847635042877,"score_spread":0.18320193575494717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W196214544","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026429897,0.00022347194,0.9650934,0.00025066486,0.00008230384,0.00007752419,0.0003417023,0.005190057,0.0023110264],"genre_scores_gemma":[0.35600737,0.00034754814,0.6329969,0.00017823947,0.00012791545,0.00022059988,0.0018801297,0.0006732316,0.0075680437],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994387,0.00021370537,0.000038324397,0.00015167672,0.00012259911,0.00003497452],"domain_scores_gemma":[0.99887437,0.0006064516,0.00009012345,0.00020968243,0.00018876685,0.000030579733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073956384,0.0007678604,0.00053745334,0.00056971534,0.00029396976,0.0007323058,0.0010768295,0.0007901358,0.0036643422],"category_scores_gemma":[0.004543199,0.00037503458,0.00070335146,0.00065884186,0.00034182044,0.0016370012,0.0008149799,0.0008455946,0.0023230773],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021569566,0.00008671484,0.0009359212,0.00023630308,0.00007503143,0.0003330244,0.00028925797,0.40166652,0.021024266,0.01586892,0.009547125,0.54972124],"study_design_scores_gemma":[0.000008880988,0.000023924877,0.000101133395,0.00000553905,0.0000087939525,0.00002877259,0.000012665345,0.987852,0.0036918956,0.006871236,0.0013882681,0.0000069797343],"about_ca_topic_score_codex":0.0020188582,"about_ca_topic_score_gemma":0.0035446908,"teacher_disagreement_score":0.0036643422,"about_ca_system_score_codex":0.000457573,"about_ca_system_score_gemma":0.00037485207,"threshold_uncertainty_score":0.01225847},"labels":[],"label_agreement":null},{"id":"W1963912989","doi":"10.1145/1236181.1236183","title":"Inferential language models for information retrieval","year":2006,"lang":"en","type":"article","venue":"ACM Transactions on Asian Language Information Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Inference; Language model; Smoothing; Term (time); Natural language processing; Artificial intelligence; Query language; Machine learning; Information retrieval; Data mining","score_opus":0.010749421293758955,"score_gpt":0.24988615467232533,"score_spread":0.2391367333785664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1963912989","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017254103,0.003045057,0.98818713,0.0012676676,0.00010571864,0.0001032351,0.00035642562,0.0005521297,0.0046571707],"genre_scores_gemma":[0.24201608,0.008206056,0.7339955,0.0012844715,0.0012571277,0.0014361087,0.0023095133,0.00033347728,0.009161644],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99334806,0.0039762533,0.00040875314,0.00085522863,0.0012252646,0.00018639292],"domain_scores_gemma":[0.9853462,0.012093166,0.0005664745,0.0010956906,0.0007728639,0.00012555478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069538113,0.0017643084,0.001620971,0.003011188,0.0011100797,0.0038692846,0.0029164904,0.002469426,0.00589927],"category_scores_gemma":[0.023158943,0.0008761336,0.0025633983,0.0035118556,0.0026599031,0.008618252,0.0020198452,0.004085399,0.0023749436],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000065841676,0.000088461195,0.0004947205,0.00029847742,0.00012761174,0.00021422897,0.00038687835,0.08311794,0.00045560751,0.85053825,0.005026866,0.05918516],"study_design_scores_gemma":[0.000018631838,0.000021411232,0.0000849099,0.00004503776,0.00003228945,0.00006863752,0.000038384656,0.28212473,0.00021429709,0.70905864,0.0082678795,0.000025212505],"about_ca_topic_score_codex":0.0058086836,"about_ca_topic_score_gemma":0.004097866,"teacher_disagreement_score":0.0069538113,"about_ca_system_score_codex":0.0029200937,"about_ca_system_score_gemma":0.0019395652,"threshold_uncertainty_score":0.03677571},"labels":[],"label_agreement":null},{"id":"W1965221819","doi":"10.1016/j.patcog.2009.06.003","title":"Personalized text snippet extraction using statistical language models","year":2009,"lang":"en","type":"article","venue":"Pattern Recognition","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"National Natural Science Foundation of China","keywords":"Snippet; Computer science; Automatic summarization; Information retrieval; Process (computing); Language model; Personalized search; Question answering; Identification (biology); Information extraction; Text graph; Task (project management); Search engine; Natural language processing; Artificial intelligence","score_opus":0.07131953149087347,"score_gpt":0.31790281021732764,"score_spread":0.24658327872645416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965221819","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062268816,0.0015709093,0.82241076,0.0009624486,0.00057227624,0.000856969,0.028981665,0.07543848,0.006937724],"genre_scores_gemma":[0.22813831,0.0016223397,0.6647262,0.00027988694,0.0006595607,0.00084378285,0.07907272,0.003635295,0.021021884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991493,0.000122050405,0.000095219715,0.00025939645,0.00029870588,0.00007529468],"domain_scores_gemma":[0.998035,0.00072479586,0.00017430006,0.00029017933,0.00067720545,0.00009847249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004987044,0.0015031612,0.0012752406,0.0067355307,0.0007246024,0.0012965057,0.00092646206,0.00095745217,0.009234997],"category_scores_gemma":[0.0031381529,0.00056407624,0.0014288492,0.00406592,0.0001957848,0.0022187373,0.00087532506,0.0010337752,0.013526282],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006499365,0.0003751293,0.004881597,0.0007505398,0.00021170209,0.00091588893,0.00023488169,0.009547928,0.06277293,0.002418986,0.07957127,0.8376692],"study_design_scores_gemma":[0.0001846981,0.0004823197,0.014343829,0.00016723025,0.00069535454,0.0023709484,0.0006193231,0.6937762,0.16006203,0.01666712,0.110448785,0.0001822019],"about_ca_topic_score_codex":0.0025379006,"about_ca_topic_score_gemma":0.0059400164,"teacher_disagreement_score":0.009234997,"about_ca_system_score_codex":0.00040766655,"about_ca_system_score_gemma":0.0010712867,"threshold_uncertainty_score":0.03089416},"labels":[],"label_agreement":null},{"id":"W1965599733","doi":"10.1016/j.mcm.2009.08.015","title":"A behavioural mode research on user-focus summarization","year":2009,"lang":"en","type":"article","venue":"Mathematical and Computer Modelling","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"","keywords":"Automatic summarization; Computer science; Multi-document summarization; Focus (optics); Selection (genetic algorithm); Information retrieval; Process (computing); Quality (philosophy); tf–idf; Granularity; Artificial intelligence; Term (time)","score_opus":0.12966549870643815,"score_gpt":0.33046489513329735,"score_spread":0.2007993964268592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965599733","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31268743,0.0009515254,0.6112114,0.0021774278,0.0001371069,0.00066846097,0.00055395084,0.0018230431,0.06978969],"genre_scores_gemma":[0.9522651,0.00022587956,0.03971513,0.00018847888,0.00007000486,0.00025691738,0.00020470157,0.00028049707,0.00679333],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9932848,0.003894714,0.00030952456,0.0010446842,0.0011472397,0.00031896],"domain_scores_gemma":[0.9422803,0.041241158,0.0031247612,0.0057144715,0.0065715667,0.0010677367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076608486,0.00067428045,0.00068836234,0.0019576857,0.0010759014,0.0043996116,0.0022356363,0.001568716,0.015191546],"category_scores_gemma":[0.058662724,0.0010434561,0.00086464407,0.0015590391,0.0016711623,0.009015967,0.0018226244,0.0016255481,0.0022576055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023104902,0.00075771357,0.060068563,0.0022392247,0.00048214998,0.00036098558,0.10479016,0.006350178,0.0877633,0.37400246,0.0071396325,0.3537351],"study_design_scores_gemma":[0.00052269927,0.0033352806,0.14402996,0.0008483486,0.0012305886,0.002861222,0.056427553,0.30767542,0.07351115,0.34097803,0.067852184,0.0007275126],"about_ca_topic_score_codex":0.0027995182,"about_ca_topic_score_gemma":0.0011573634,"teacher_disagreement_score":0.015191546,"about_ca_system_score_codex":0.0014356156,"about_ca_system_score_gemma":0.0008558007,"threshold_uncertainty_score":0.050820827},"labels":[],"label_agreement":null},{"id":"W1965732604","doi":"10.1109/iccse.2014.6926421","title":"Mining case summaries in BioWorld","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Correctness; Information retrieval; Natural language processing; Artificial intelligence; Data science; Data mining; Programming language","score_opus":0.028524374019517164,"score_gpt":0.24558018467743797,"score_spread":0.21705581065792082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965732604","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70275563,0.005331938,0.2561709,0.002457237,0.00038553847,0.0014859149,0.022263365,0.003608579,0.0055409027],"genre_scores_gemma":[0.70799446,0.0016407454,0.25411242,0.0002004845,0.00022442921,0.00063721434,0.03341827,0.00013664385,0.0016353184],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968464,0.0010904662,0.000673801,0.00057043735,0.0007203164,0.000098529454],"domain_scores_gemma":[0.9650814,0.024573613,0.0043789227,0.002215624,0.00304223,0.0007082536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022744006,0.00080881425,0.00056037284,0.008686253,0.00079502544,0.0018866019,0.0009201071,0.000985959,0.0017680739],"category_scores_gemma":[0.03530508,0.00026342814,0.00055575056,0.0055189566,0.000396765,0.0023511436,0.0009740698,0.00059609464,0.00048562474],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010133279,0.0011252668,0.117066264,0.0029901157,0.00053794374,0.015439614,0.007397528,0.04418083,0.01888472,0.01023223,0.04041179,0.7407204],"study_design_scores_gemma":[0.00037069802,0.0013754666,0.09949397,0.0015111351,0.0008294421,0.0149410395,0.016666226,0.54538816,0.06257669,0.06797876,0.18855229,0.00031608483],"about_ca_topic_score_codex":0.0012232442,"about_ca_topic_score_gemma":0.002449714,"teacher_disagreement_score":0.008686253,"about_ca_system_score_codex":0.0006712855,"about_ca_system_score_gemma":0.00090538355,"threshold_uncertainty_score":0.012028277},"labels":[],"label_agreement":null},{"id":"W1967932323","doi":"10.1145/2396761.2398483","title":"Hierarchical topic integration through semi-supervised hierarchical topic modeling","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Natural Science Foundation of China","keywords":"Computer science; Topic model; Hierarchical database model; Hierarchical organization; Information retrieval; Data science; Artificial intelligence; Data mining","score_opus":0.05065778021385397,"score_gpt":0.2817311694565262,"score_spread":0.23107338924267223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967932323","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008593137,0.00021957423,0.9893741,0.00007699912,0.000018321378,0.00007513513,0.000102152204,0.0011990295,0.00034160216],"genre_scores_gemma":[0.27894372,0.00040998793,0.71447384,0.00015045734,0.00022107553,0.00046589802,0.0023538866,0.00048461996,0.0024965208],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966983,0.0012751338,0.00020950916,0.00090145256,0.00073156366,0.00018407976],"domain_scores_gemma":[0.9938048,0.0035668833,0.0005970834,0.00078399177,0.0010619878,0.00018522024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038183029,0.0015723514,0.0016849699,0.0033369039,0.000851194,0.0017497529,0.0022730734,0.0012513335,0.0014052318],"category_scores_gemma":[0.0093398085,0.00096398767,0.0020945838,0.0029711355,0.0007777309,0.0041628745,0.002193453,0.0023244647,0.0012161267],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006531103,0.0007708779,0.0070787547,0.00052175677,0.0006988072,0.00033731194,0.001957093,0.23806955,0.025861563,0.020071507,0.012451269,0.69152844],"study_design_scores_gemma":[0.00002299172,0.000044121687,0.00066236046,0.00001438512,0.00006676601,0.00006719444,0.00006291562,0.98366654,0.0034029584,0.010268183,0.0016961951,0.000025411146],"about_ca_topic_score_codex":0.0036843433,"about_ca_topic_score_gemma":0.006220844,"teacher_disagreement_score":0.0038183029,"about_ca_system_score_codex":0.000801598,"about_ca_system_score_gemma":0.0014174667,"threshold_uncertainty_score":0.020193338},"labels":[],"label_agreement":null},{"id":"W1968377807","doi":"10.1080/13645579.2011.645700","title":"A computer-assisted approach to filtering large numbers of documents for media analyses","year":2012,"lang":"en","type":"article","venue":"International Journal of Social Research Methodology","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Trinity Western University; Western University; Abbotsford Veterinary Clinic","funders":"Trinity Western University","keywords":"Computer science; Selection (genetic algorithm); Filter (signal processing); Selection bias; Information retrieval; Reduction (mathematics); Data mining; Data science; Machine learning; Statistics; Mathematics","score_opus":0.6865705192359098,"score_gpt":0.610141181123146,"score_spread":0.07642933811276387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1968377807","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0150260655,0.0005439094,0.9449423,0.00044750026,0.0001867283,0.002113743,0.00194281,0.032908272,0.0018886592],"genre_scores_gemma":[0.037491996,0.00018946639,0.9567405,0.00017270188,0.00026787064,0.0013100869,0.0013991662,0.00043247794,0.0019957598],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99018854,0.002941701,0.0014778005,0.0020287372,0.0030820516,0.0002811368],"domain_scores_gemma":[0.93754584,0.042483408,0.0033932962,0.005647536,0.010105562,0.0008243469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014827678,0.001727318,0.0023972313,0.01781333,0.002471493,0.0043024183,0.0020270196,0.0020920786,0.010507875],"category_scores_gemma":[0.042141907,0.0009577386,0.0014606655,0.012428109,0.00079081947,0.0035293752,0.0027054597,0.0016542044,0.006859697],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006996467,0.0005425259,0.0057498426,0.0007726457,0.00040064438,0.00030322955,0.0015792237,0.0027786146,0.04775783,0.0028000136,0.022218803,0.91439706],"study_design_scores_gemma":[0.0009224784,0.0013239745,0.032744803,0.0004611521,0.0012028053,0.003264823,0.0022252325,0.6034699,0.13296396,0.03455442,0.18608293,0.0007836027],"about_ca_topic_score_codex":0.0038044997,"about_ca_topic_score_gemma":0.0068576504,"teacher_disagreement_score":0.01781333,"about_ca_system_score_codex":0.0007218812,"about_ca_system_score_gemma":0.0025657022,"threshold_uncertainty_score":0.07841724},"labels":[],"label_agreement":null},{"id":"W1970791442","doi":"10.1109/asru.2011.6163952","title":"A dialogue system for accessing drug reviews","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Grassroots; Parsing; Domain (mathematical analysis); Recommender system; World Wide Web; Drug; Information retrieval; Natural language processing; Medicine","score_opus":0.13717618333409023,"score_gpt":0.2916642548653229,"score_spread":0.15448807153123265,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1970791442","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049876023,0.0010003448,0.8309,0.0006069545,0.00033594386,0.0006834864,0.0036548004,0.106887035,0.0060554463],"genre_scores_gemma":[0.38659644,0.00039938206,0.5902539,0.0006067513,0.00034360398,0.0010272041,0.0086515015,0.0017695762,0.010351643],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99882156,0.0005185898,0.000099908044,0.0002850521,0.00022910022,0.000045761306],"domain_scores_gemma":[0.9960974,0.0026461545,0.00017184396,0.0002683767,0.0005793554,0.00023694868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018062976,0.00077711936,0.00081262406,0.0009709584,0.0005781088,0.0008981824,0.0009748747,0.0010450762,0.0075001316],"category_scores_gemma":[0.006340267,0.00035723328,0.00048321602,0.0004084926,0.0002393078,0.0013705493,0.00113155,0.0007984891,0.0047737337],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019227145,0.0006093332,0.0053004962,0.0016233468,0.00025476993,0.0010179617,0.0026036962,0.0059402413,0.1402306,0.005295017,0.057958703,0.77724314],"study_design_scores_gemma":[0.00073305645,0.001831419,0.015712723,0.00024843754,0.0006309375,0.0032931764,0.0011929114,0.55710834,0.15356031,0.012502165,0.25269914,0.00048734527],"about_ca_topic_score_codex":0.0015801168,"about_ca_topic_score_gemma":0.0014167198,"teacher_disagreement_score":0.0075001316,"about_ca_system_score_codex":0.0003177132,"about_ca_system_score_gemma":0.00056936196,"threshold_uncertainty_score":0.025090456},"labels":[],"label_agreement":null},{"id":"W1974230338","doi":"10.1109/ccece.2013.6567775","title":"Virtual cardiologist &amp;#x2014; A conversational system for medical diagnosis","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Conversation; Computer science; Focus (optics); Meaning (existential); Process (computing); Medical diagnosis; Ask price; Human–computer interaction; Multimedia; Psychology; Programming language; Medicine; Radiology","score_opus":0.03184602320474247,"score_gpt":0.2570850538251836,"score_spread":0.22523903062044112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974230338","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0772756,0.00097456464,0.83957833,0.0029327655,0.00032506918,0.00048195242,0.0024745588,0.053262934,0.022694368],"genre_scores_gemma":[0.6597473,0.0002943109,0.3232784,0.0007820922,0.00022665459,0.0004270001,0.0028280644,0.0007760521,0.011640028],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989403,0.00062136434,0.00004647346,0.00019943935,0.00013069129,0.000061860905],"domain_scores_gemma":[0.9976513,0.0015991383,0.00008000481,0.00019470569,0.00018796157,0.00028699503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021213717,0.00047052023,0.00032078478,0.0006050931,0.000810185,0.001067997,0.0010592559,0.000983062,0.013812359],"category_scores_gemma":[0.0060813627,0.0003156204,0.0003201038,0.00027210635,0.0002643258,0.0013059045,0.0015947805,0.00085162726,0.0040300405],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026056347,0.0006003805,0.00852738,0.00059220265,0.000169276,0.0014665993,0.0062929904,0.015493469,0.061714955,0.019280696,0.09874933,0.78450716],"study_design_scores_gemma":[0.00045800468,0.00086486456,0.006418335,0.00015410843,0.00024363375,0.0018976941,0.0016410172,0.7236632,0.036715567,0.026674828,0.20102727,0.00024145002],"about_ca_topic_score_codex":0.002858061,"about_ca_topic_score_gemma":0.0023470873,"teacher_disagreement_score":0.013812359,"about_ca_system_score_codex":0.00050011603,"about_ca_system_score_gemma":0.0009794091,"threshold_uncertainty_score":0.04620695},"labels":[],"label_agreement":null},{"id":"W1974442733","doi":"10.5539/ijel.v2n3p64","title":"Metadiscoursal Markers in Medical and Literary Texts","year":2012,"lang":"en","type":"article","venue":"International Journal of English Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Linguistics; Test (biology); Significant difference; Medical literature; Literature; Psychology; Statistics; Mathematics; Biology; Art; Medicine; Philosophy; Pathology; Botany","score_opus":0.0140155748697388,"score_gpt":0.2926779781701875,"score_spread":0.27866240330044867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974442733","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9800463,0.0040366882,0.005933728,0.00037805928,0.00008063931,0.00013336318,0.00046653385,0.00006316842,0.008861637],"genre_scores_gemma":[0.99132687,0.000887673,0.0062786173,0.000033041397,0.000053064476,0.0001018239,0.00033825118,0.000019809673,0.00096083456],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9957092,0.0019240294,0.00084908836,0.00042128837,0.0009745334,0.00012190597],"domain_scores_gemma":[0.9605202,0.028484527,0.005654277,0.0015774257,0.0029184592,0.000845121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002727158,0.00025037242,0.00042859378,0.013492804,0.00132378,0.0031567316,0.0004195061,0.0004547663,0.0021851344],"category_scores_gemma":[0.025989544,0.00023704035,0.00026528593,0.010080276,0.0016471125,0.0027907437,0.0019967325,0.00046520043,0.000299439],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025664063,0.0004133811,0.17852421,0.005619303,0.00023537847,0.0043768766,0.40200582,0.0006952236,0.042249825,0.044734567,0.0023848864,0.31619412],"study_design_scores_gemma":[0.00010964691,0.0008310887,0.5761832,0.0026631677,0.00041560142,0.010984237,0.23174803,0.0043991706,0.021066813,0.018154513,0.1332266,0.00021792983],"about_ca_topic_score_codex":0.00049528736,"about_ca_topic_score_gemma":0.00073876645,"teacher_disagreement_score":0.013492804,"about_ca_system_score_codex":0.0010853285,"about_ca_system_score_gemma":0.0007212768,"threshold_uncertainty_score":0.014422715},"labels":[],"label_agreement":null},{"id":"W1974586764","doi":"10.1080/09540090110052996","title":"Syntactic systematicity arising from semantic predictions in a Hebbian-competitive network","year":2001,"lang":"en","type":"article","venue":"Connection Science","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Hebbian theory; Computer science; Artificial intelligence; Natural language processing; Connectionism; Embedding; Competitive learning; Simple (philosophy); Leabra; Artificial neural network; Generalization error; Philosophy","score_opus":0.031078408995862356,"score_gpt":0.2704708189559582,"score_spread":0.23939240996009584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974586764","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30742052,0.00035490646,0.67976743,0.0007819784,0.000045028875,0.000055615703,0.00011170142,0.00045294425,0.011009989],"genre_scores_gemma":[0.97142166,0.00011221071,0.026040325,0.00004845039,0.000018039269,0.000028515648,0.00007099836,0.000024601844,0.0022351358],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99975365,0.00009045682,0.000010741568,0.00007192965,0.000042782194,0.000030494071],"domain_scores_gemma":[0.99817324,0.0011573522,0.00018197308,0.0001817594,0.00022264001,0.00008309434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00090767466,0.0002726889,0.00031770108,0.00034897303,0.00038656738,0.0006984164,0.00084064546,0.00091237074,0.0014396124],"category_scores_gemma":[0.00497389,0.00033573274,0.00029920074,0.00039644612,0.0010670003,0.0017520789,0.0005951792,0.0007623364,0.00023831414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031269778,0.00009194236,0.0074409605,0.00020064418,0.00016490971,0.00043976714,0.00076002814,0.5116457,0.031359658,0.31625667,0.0029005876,0.12842649],"study_design_scores_gemma":[0.000007635592,0.000032921886,0.0005485168,0.0000050586773,0.000014533238,0.00003884821,0.00001625458,0.89886403,0.0012888241,0.09865675,0.0005187456,0.000007930883],"about_ca_topic_score_codex":0.002736367,"about_ca_topic_score_gemma":0.0047081737,"teacher_disagreement_score":0.002736367,"about_ca_system_score_codex":0.000770255,"about_ca_system_score_gemma":0.0007561636,"threshold_uncertainty_score":0.0055886507},"labels":[],"label_agreement":null},{"id":"W1975422446","doi":"10.1145/1076034.1076086","title":"Integrating word relationships into language models","year":2005,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":166,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"WordNet; Computer science; Dependency (UML); Language model; Natural language processing; Word (group theory); Artificial intelligence; Independence (probability theory); Information retrieval; Linguistics","score_opus":0.045667234214624874,"score_gpt":0.2723542965639025,"score_spread":0.2266870623492776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975422446","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071467967,0.00092907035,0.986783,0.0005218282,0.00008070186,0.00010135428,0.00034001048,0.0012067353,0.0028905687],"genre_scores_gemma":[0.3459547,0.0028484713,0.63817865,0.0008790731,0.0004206543,0.0006285964,0.0024581202,0.0007921054,0.007839608],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980167,0.0009410016,0.00015059298,0.00039113912,0.00040772287,0.000092814815],"domain_scores_gemma":[0.9964586,0.0024614965,0.0002590692,0.00035553938,0.00039479343,0.00007047885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022381938,0.0011993148,0.0009403174,0.0023745087,0.00060488837,0.0021165763,0.0016052992,0.0010790706,0.002685411],"category_scores_gemma":[0.00928036,0.0008240663,0.0020310953,0.0021821547,0.00056928524,0.007745254,0.0015142665,0.0018527886,0.0026085975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024088802,0.00030011812,0.0052234195,0.0008486022,0.00095039123,0.00051511626,0.001156631,0.34266436,0.013498138,0.16656472,0.012608812,0.45542887],"study_design_scores_gemma":[0.000024146862,0.00009218153,0.0007188988,0.00006187382,0.0002769974,0.00025294343,0.00011110777,0.86297286,0.002598979,0.11030586,0.022511264,0.00007284874],"about_ca_topic_score_codex":0.0048170527,"about_ca_topic_score_gemma":0.0071948543,"teacher_disagreement_score":0.0048170527,"about_ca_system_score_codex":0.00088013103,"about_ca_system_score_gemma":0.0013552997,"threshold_uncertainty_score":0.011836886},"labels":[],"label_agreement":null},{"id":"W1975471657","doi":"10.3758/bf03201255","title":"Turning an advantage into a disadvantage: Ambiguity effects in lexical decision versus reading tasks","year":2000,"lang":"en","type":"article","venue":"Memory & Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"","keywords":"Lexical decision task; Categorization; Psychology; Ambiguity; Reading (process); Task (project management); Contrast (vision); Meaning (existential); Linguistics; Word (group theory); Cognitive psychology; Lexico; Disadvantage; Word recognition; Natural language processing; Artificial intelligence; Cognition; Computer science","score_opus":0.018558864728846963,"score_gpt":0.3081781869163606,"score_spread":0.28961932218751363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1975471657","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9919844,0.00036156256,0.0019523252,0.00037164145,0.00008456722,0.000026976795,0.000112488364,0.000076554075,0.0050294623],"genre_scores_gemma":[0.9960522,0.00014336091,0.0018208688,0.000331023,0.00014645945,0.000040799183,0.0001822248,0.00028728798,0.0009959105],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9972451,0.0007288652,0.00030584537,0.00065752823,0.0008291225,0.000233537],"domain_scores_gemma":[0.92329484,0.06618809,0.0039201267,0.0033868672,0.0014178073,0.0017922166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005987213,0.0008352992,0.0015332244,0.0009822708,0.0005221484,0.0037784278,0.0009975025,0.002580806,0.007392044],"category_scores_gemma":[0.08006375,0.00084334635,0.0005348412,0.0007382745,0.0010653493,0.007589364,0.0023871765,0.0027789525,0.0013454433],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.033149727,0.004073207,0.099849835,0.0017504082,0.0006249033,0.0013331296,0.01447835,0.0028885426,0.65919197,0.010991483,0.0046397992,0.16702868],"study_design_scores_gemma":[0.002805766,0.0062890938,0.76560354,0.0002838291,0.0013774561,0.0031206456,0.0060163406,0.026511742,0.06890072,0.11408627,0.00445115,0.0005534442],"about_ca_topic_score_codex":0.00084828166,"about_ca_topic_score_gemma":0.0008011494,"teacher_disagreement_score":0.007392044,"about_ca_system_score_codex":0.00025657724,"about_ca_system_score_gemma":0.00038408296,"threshold_uncertainty_score":0.031663775},"labels":[],"label_agreement":null},{"id":"W1976320987","doi":"10.1142/s0218213008003881","title":"TEXT SUMMARIZATION USING LEXICAL COHESION: APPROACHES AND EVALUATIONS","year":2008,"lang":"en","type":"article","venue":"International Journal of Artificial Intelligence Tools","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Automatic summarization; Cohesion (chemistry); Popularity; Multi-document summarization; Information retrieval; Context (archaeology); Identification (biology); The Internet; World Wide Web; Natural language processing","score_opus":0.3494098354123179,"score_gpt":0.3781528857572114,"score_spread":0.028743050344893495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1976320987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.832915,0.017926835,0.11474091,0.0007424783,0.00026390603,0.0035457453,0.0027383657,0.010476136,0.016650578],"genre_scores_gemma":[0.791146,0.005095046,0.18906718,0.00013199534,0.00023432808,0.0011878584,0.00861583,0.00047179317,0.0040499824],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99079376,0.005179578,0.0009154787,0.00070783775,0.0021484995,0.00025481734],"domain_scores_gemma":[0.97590977,0.018717399,0.0008845854,0.001073281,0.002896041,0.00051896693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007890316,0.0013272447,0.0015799257,0.004940256,0.0010772337,0.0021264309,0.0014976755,0.0013515276,0.0017534222],"category_scores_gemma":[0.02345077,0.00035035057,0.0007332217,0.0044062063,0.0007008998,0.0035029692,0.0014013267,0.00079183216,0.00063205406],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031328776,0.0028137236,0.009672636,0.00534467,0.00081679906,0.00039766286,0.003790224,0.025685303,0.032624777,0.0019635537,0.010740754,0.903017],"study_design_scores_gemma":[0.0027240936,0.016497875,0.0784461,0.0006395018,0.0037314175,0.0013520925,0.010703406,0.71281034,0.123046406,0.006382276,0.043187335,0.0004791644],"about_ca_topic_score_codex":0.008003002,"about_ca_topic_score_gemma":0.008066168,"teacher_disagreement_score":0.008003002,"about_ca_system_score_codex":0.0012509205,"about_ca_system_score_gemma":0.000953374,"threshold_uncertainty_score":0.041728497},"labels":[],"label_agreement":null},{"id":"W1977123279","doi":"10.1080/08839510903078093","title":"QUESTION ANSWERING USING QUESTION CLASSIFICATION AND DOCUMENT TAGGING","year":2009,"lang":"en","type":"article","venue":"Applied Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Question answering; Computer science; Information retrieval; Document retrieval; Document classification; Natural language processing; Artificial intelligence","score_opus":0.06235094333815122,"score_gpt":0.3230424587390867,"score_spread":0.2606915154009355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1977123279","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01325611,0.00080350396,0.9700169,0.0006111961,0.00017248212,0.0009009062,0.0011825283,0.0075109247,0.005545545],"genre_scores_gemma":[0.12238787,0.0005987462,0.8638506,0.00048208344,0.00028038313,0.0007917214,0.0065891165,0.00044769474,0.004571813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9841875,0.008795626,0.0014793781,0.002497297,0.002463596,0.0005766855],"domain_scores_gemma":[0.9657149,0.0221302,0.001765636,0.0041056965,0.005731643,0.00055194914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011562699,0.0015883559,0.0017519488,0.01016402,0.0015764326,0.0044760993,0.0023560512,0.0024187844,0.0043967133],"category_scores_gemma":[0.0278323,0.000605192,0.0019988804,0.006833223,0.0011217529,0.006111993,0.0025706394,0.0019268666,0.0044745207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040962436,0.0008747241,0.007766747,0.00095840055,0.0002484726,0.00023718514,0.0021999907,0.009939582,0.025365299,0.020216297,0.027800817,0.90398294],"study_design_scores_gemma":[0.00022652418,0.0005609731,0.011285313,0.0005157277,0.00048712722,0.0011766976,0.0015987451,0.6890668,0.07258618,0.11267301,0.109455,0.00036792638],"about_ca_topic_score_codex":0.005982258,"about_ca_topic_score_gemma":0.004957814,"teacher_disagreement_score":0.011562699,"about_ca_system_score_codex":0.0016553978,"about_ca_system_score_gemma":0.00163819,"threshold_uncertainty_score":0.061150134},"labels":[],"label_agreement":null},{"id":"W1979325497","doi":"10.4018/jswis.2006070104","title":"Information Retrieval by Semantic Similarity","year":2006,"lang":"en","type":"article","venue":"International Journal on Semantic Web and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":228,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Semantic similarity; Computer science; Information retrieval; WordNet; Explicit semantic analysis; Semantic integration; Semantic computing; Similarity (geometry); Ontology; Vector space model; Semantic search; Natural language processing; Artificial intelligence; Semantic technology; Semantic Web; Image (mathematics)","score_opus":0.008796478971780134,"score_gpt":0.23001016622555082,"score_spread":0.2212136872537707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979325497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010997064,0.016318308,0.9412794,0.0025766124,0.0007587394,0.0014659702,0.0014674026,0.0019037507,0.023232665],"genre_scores_gemma":[0.18804897,0.020510517,0.7680677,0.001323306,0.0017617714,0.0018579127,0.0059331344,0.00032163112,0.012175075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9868699,0.005352171,0.0013224863,0.0015695249,0.004513247,0.0003727209],"domain_scores_gemma":[0.99501294,0.0023993077,0.00038326255,0.0013081385,0.0007929443,0.00010338329],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006391887,0.001549492,0.0033272663,0.016642332,0.0013348731,0.0074669276,0.0025669595,0.0029509864,0.007896899],"category_scores_gemma":[0.02170737,0.0006487645,0.002246656,0.021507237,0.002684258,0.017833991,0.0055820397,0.002043382,0.006432825],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022674364,0.00030504068,0.0013640956,0.0023954783,0.00051916734,0.00033629726,0.0006453901,0.015523001,0.007902218,0.27907553,0.036553096,0.6551539],"study_design_scores_gemma":[0.00014983948,0.00041115357,0.0014710529,0.000590718,0.00026758763,0.0013962333,0.000719341,0.15196039,0.009958476,0.6972605,0.13558356,0.00023106778],"about_ca_topic_score_codex":0.0016393336,"about_ca_topic_score_gemma":0.0012089687,"teacher_disagreement_score":0.016642332,"about_ca_system_score_codex":0.0025845733,"about_ca_system_score_gemma":0.0020599167,"threshold_uncertainty_score":0.03380394},"labels":[],"label_agreement":null},{"id":"W1979668205","doi":"10.1109/ictai.2013.108","title":"Assessing Procedural Knowledge in Free-Text Answers through a Hybrid Semantic Web Approach","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Novelty; Text messaging; Natural language processing; Semantic similarity; Grading (engineering); Artificial intelligence; Information retrieval; WordNet; World Wide Web","score_opus":0.036868468348414875,"score_gpt":0.27796125148350354,"score_spread":0.24109278313508867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1979668205","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15998404,0.00032139017,0.83027667,0.00016043914,0.000035314253,0.00036202857,0.00044813196,0.0016402823,0.0067717],"genre_scores_gemma":[0.74310946,0.00017409369,0.25252852,0.000046256402,0.000035577195,0.00031100176,0.00073621835,0.00009171451,0.0029671467],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99530894,0.0016378324,0.00036287447,0.0005859427,0.0019405275,0.00016390347],"domain_scores_gemma":[0.9923598,0.0037473072,0.00086981343,0.0005848102,0.002159642,0.0002786137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037285907,0.00087888696,0.0007403691,0.011616586,0.00051752274,0.0026954557,0.0010935997,0.0012897729,0.0022285127],"category_scores_gemma":[0.0099609215,0.00029722674,0.0009552526,0.003907287,0.0006457376,0.0046885023,0.002147928,0.00059333816,0.0008353398],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009440799,0.0009795391,0.04088592,0.0006716303,0.00048409033,0.00024984786,0.0020324066,0.027362196,0.033552825,0.016654402,0.0018767975,0.8743062],"study_design_scores_gemma":[0.00012131259,0.0007377306,0.067771316,0.00016052296,0.000384592,0.0006479655,0.0025413432,0.81013787,0.048380345,0.057435468,0.011480668,0.00020091893],"about_ca_topic_score_codex":0.0017282483,"about_ca_topic_score_gemma":0.0027500833,"teacher_disagreement_score":0.011616586,"about_ca_system_score_codex":0.00073991803,"about_ca_system_score_gemma":0.0008089022,"threshold_uncertainty_score":0.019718885},"labels":[],"label_agreement":null},{"id":"W1982323159","doi":"10.3115/1610075.1610081","title":"Distributional measures of concept-distance","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; University of Toronto","keywords":"WordNet; Computer science; Distance measures; Word (group theory); Natural language processing; Ranking (information retrieval); Task (project management); Artificial intelligence; Distance matrix; Semantic similarity; Distance measurement; Distributional semantics; Thesaurus; Information retrieval; Mathematics; Algorithm","score_opus":0.0293469860076733,"score_gpt":0.23340295520841797,"score_spread":0.20405596920074467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982323159","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029675594,0.0017135193,0.9609545,0.00037963587,0.000121303514,0.00011181154,0.0004414679,0.00027957564,0.0063225203],"genre_scores_gemma":[0.549742,0.000907831,0.4455555,0.00017314049,0.00030629142,0.0005368026,0.0010338214,0.00013687766,0.001607668],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9938636,0.0020426242,0.0006864082,0.0010922228,0.0021232972,0.00019189143],"domain_scores_gemma":[0.9812656,0.010032496,0.0022160436,0.0029109211,0.0029732224,0.0006017613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005197365,0.0007693791,0.0009011338,0.010013173,0.0011411231,0.0033567653,0.0020787467,0.0010744691,0.0029026389],"category_scores_gemma":[0.035382092,0.00040405273,0.0010853807,0.007397589,0.002323378,0.008225145,0.0036216422,0.0017033053,0.0008109709],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012452081,0.00015927183,0.009231015,0.00044921774,0.000296679,0.00015245483,0.00091459334,0.025657084,0.0037597888,0.6928276,0.002581341,0.26384652],"study_design_scores_gemma":[0.000025647007,0.000180688,0.005643195,0.00010201962,0.00007103931,0.00060402934,0.00038849172,0.11874701,0.002993345,0.85587937,0.015263664,0.00010153838],"about_ca_topic_score_codex":0.0009250869,"about_ca_topic_score_gemma":0.00089482067,"teacher_disagreement_score":0.010013173,"about_ca_system_score_codex":0.0015050969,"about_ca_system_score_gemma":0.0010368159,"threshold_uncertainty_score":0.027486622},"labels":[],"label_agreement":null},{"id":"W1982795102","doi":"10.1109/icalt.2010.47","title":"Automarking: Automatic Assessment of Open Questions","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"Athabasca University","keywords":"Computer science; Meaning (existential); Leverage (statistics); Question answering; Probabilistic logic; Taxonomy (biology); Natural language processing; Artificial intelligence; Psychology","score_opus":0.03326472433835425,"score_gpt":0.3561803345393306,"score_spread":0.32291561020097637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982795102","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057379745,0.0013485289,0.84871393,0.00066988263,0.00034164713,0.0022177119,0.0052001597,0.07744686,0.00668146],"genre_scores_gemma":[0.26495874,0.0006624977,0.7073035,0.0003125804,0.00034109046,0.0016867564,0.014306939,0.0014515893,0.00897628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99085367,0.0042508557,0.0007626141,0.0016867858,0.0021200397,0.0003260579],"domain_scores_gemma":[0.96840745,0.021544198,0.0018362172,0.00228712,0.005173774,0.0007511238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008147277,0.0015884337,0.0012072007,0.006237664,0.00077780656,0.00319806,0.0021763667,0.0020563782,0.009720596],"category_scores_gemma":[0.036848467,0.0005601963,0.0007909367,0.0026072804,0.00062635535,0.0059459885,0.0038085936,0.0019321835,0.008629402],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060911477,0.000342371,0.0061943554,0.0006091258,0.000065230604,0.0001676297,0.0014227369,0.0016073409,0.024655974,0.003451787,0.02406722,0.9368071],"study_design_scores_gemma":[0.0004722488,0.0010754039,0.049159855,0.0005337529,0.00021536609,0.0015370391,0.0042373254,0.5187938,0.1569704,0.07466489,0.1919267,0.00041323382],"about_ca_topic_score_codex":0.0019281582,"about_ca_topic_score_gemma":0.0021245629,"teacher_disagreement_score":0.009720596,"about_ca_system_score_codex":0.0008689463,"about_ca_system_score_gemma":0.0015253332,"threshold_uncertainty_score":0.043087482},"labels":[],"label_agreement":null},{"id":"W198356813","doi":"","title":"Unsupervised Approach for Selecting Sentences in Query-based Summarization","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval; Relevance (law); Cosine similarity; Multi-document summarization; Task (project management); Natural language processing; Query expansion; Artificial intelligence; Cluster analysis","score_opus":0.05152992492658181,"score_gpt":0.2434124001038703,"score_spread":0.1918824751772885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W198356813","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057875473,0.0007372299,0.9349135,0.00013173734,0.00005926179,0.00039774098,0.0006742112,0.0041648187,0.0010460786],"genre_scores_gemma":[0.32763714,0.00028546693,0.6651465,0.00011624269,0.00014847644,0.00058994413,0.0034003633,0.00035163484,0.0023242864],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985353,0.0005941027,0.00014487203,0.00036161699,0.00027052237,0.000093611096],"domain_scores_gemma":[0.99753153,0.001094958,0.00020521099,0.0002634446,0.0008415712,0.000063277796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001577625,0.0011431446,0.0011248876,0.0021481009,0.0005932375,0.0007173308,0.00127275,0.0008448344,0.0011697587],"category_scores_gemma":[0.005141445,0.0003453393,0.00083725265,0.001541954,0.00034786246,0.0011777562,0.00062210154,0.00064046745,0.0010116571],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082999683,0.0004643201,0.004310172,0.00059158367,0.00029634326,0.00033748162,0.0009519727,0.057009842,0.15545347,0.004135201,0.008306525,0.76731324],"study_design_scores_gemma":[0.00011080167,0.00056632346,0.0062731085,0.000027372738,0.00024748428,0.00041631647,0.00043935084,0.89000344,0.085644305,0.0061870976,0.0099977385,0.00008669505],"about_ca_topic_score_codex":0.0019797327,"about_ca_topic_score_gemma":0.00427016,"teacher_disagreement_score":0.0021481009,"about_ca_system_score_codex":0.00041482376,"about_ca_system_score_gemma":0.00076586823,"threshold_uncertainty_score":0.008343399},"labels":[],"label_agreement":null},{"id":"W1985620863","doi":"10.3758/s13423-014-0701-7","title":"Item-properties may influence item–item associations in serial recall","year":2014,"lang":"en","type":"article","venue":"Psychonomic Bulletin & Review","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Psychology; Recall; Serial position effect; Cognitive psychology; Short Forms; Developmental psychology; Free recall; Clinical psychology","score_opus":0.030978897327540226,"score_gpt":0.26400753132175536,"score_spread":0.23302863399421514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985620863","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86388254,0.035491664,0.08782786,0.0018859832,0.00046827336,0.00040268834,0.0011562344,0.0005013387,0.008383419],"genre_scores_gemma":[0.9815533,0.004159223,0.011676493,0.0003325465,0.00027499002,0.00019589983,0.0010433081,0.00013513447,0.0006290287],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99103034,0.0043632723,0.0013902154,0.0012785122,0.0017165642,0.00022097492],"domain_scores_gemma":[0.7575281,0.20055257,0.014060817,0.01482827,0.012245743,0.0007844887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03726183,0.0008568085,0.0016186885,0.0024397888,0.00060740375,0.0046514547,0.0015246028,0.0015610995,0.003992866],"category_scores_gemma":[0.17886823,0.0012148081,0.0024649384,0.0028906171,0.0011733348,0.0053519174,0.0011443824,0.0023504002,0.0010420182],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014639946,0.000403518,0.7648804,0.0028442936,0.00668711,0.000239543,0.001902023,0.002901712,0.0065202652,0.004589873,0.002172968,0.20539436],"study_design_scores_gemma":[0.0002490535,0.00071443856,0.9335217,0.0007117466,0.005471658,0.0008824301,0.0006795683,0.012965367,0.0071061933,0.03317466,0.004391945,0.00013127353],"about_ca_topic_score_codex":0.0009452654,"about_ca_topic_score_gemma":0.0014065739,"teacher_disagreement_score":0.03726183,"about_ca_system_score_codex":0.0008439674,"about_ca_system_score_gemma":0.0009062552,"threshold_uncertainty_score":0.19706172},"labels":[],"label_agreement":null},{"id":"W1986408257","doi":"10.1145/1321440.1321453","title":"Semantic verification in an online fact seeking environment","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Question answering; Artificial intelligence; Semantic Web; Information retrieval; Semantics (computer science); Natural language processing; Word (group theory); World Wide Web; Programming language","score_opus":0.04919431739901922,"score_gpt":0.26896918553973315,"score_spread":0.21977486814071392,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986408257","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06921061,0.00016759768,0.8991947,0.0007497416,0.000056210298,0.00021378539,0.0009113541,0.026897334,0.002598698],"genre_scores_gemma":[0.43488085,0.0000913792,0.5588612,0.00025995306,0.00005606157,0.00012978268,0.0027256776,0.00085092906,0.002144259],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912565,0.003458645,0.00061944657,0.0017867296,0.0024268439,0.00045197873],"domain_scores_gemma":[0.96335864,0.024894254,0.0017866336,0.005870139,0.0034592876,0.0006311115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008789352,0.00091531023,0.0013162827,0.0032142117,0.0017822406,0.00450093,0.0024429765,0.0032117576,0.004390975],"category_scores_gemma":[0.034912027,0.00095287105,0.0019173467,0.0016421817,0.0015576688,0.010525173,0.004798495,0.0018183594,0.003005073],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025207638,0.0011042042,0.022531904,0.0008225775,0.00034801022,0.0025783882,0.0068416856,0.108489946,0.08033836,0.13188227,0.020999089,0.6215428],"study_design_scores_gemma":[0.00007666752,0.000105255815,0.0024322972,0.00006645223,0.0000823151,0.0005166973,0.00068913505,0.8706871,0.038417876,0.07042404,0.016413582,0.00008849028],"about_ca_topic_score_codex":0.007412119,"about_ca_topic_score_gemma":0.0056385146,"teacher_disagreement_score":0.008789352,"about_ca_system_score_codex":0.0014934764,"about_ca_system_score_gemma":0.0023764852,"threshold_uncertainty_score":0.04648304},"labels":[],"label_agreement":null},{"id":"W1988047293","doi":"10.1007/s10115-013-0617-y","title":"Analyzing topics and authors in chat logs for crime investigation","year":2013,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Fonds Québécois de la Recherche sur la Nature et les Technologies","keywords":"Computer science; Exploit; Latent Dirichlet allocation; Online chat; Process (computing); The Internet; Order (exchange); Volume (thermodynamics); Channel (broadcasting); Topic model; Data science; World Wide Web; Information retrieval; Computer security","score_opus":0.026709074949374544,"score_gpt":0.2508178429459265,"score_spread":0.22410876799655197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988047293","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96769696,0.00082205096,0.021146316,0.0004883923,0.000097707205,0.00015521313,0.006482151,0.00092259096,0.002188618],"genre_scores_gemma":[0.9831352,0.00027924287,0.00941956,0.00003730364,0.000111577,0.00014878472,0.0055008344,0.00007692893,0.0012905842],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99664754,0.0015421531,0.00030119545,0.000579159,0.0006717679,0.00025814897],"domain_scores_gemma":[0.9528159,0.037120104,0.0033767126,0.0021410005,0.0031849442,0.0013613895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033713803,0.00057065417,0.0006846928,0.00853906,0.0012344786,0.0022894735,0.0006488718,0.0010902622,0.0013745897],"category_scores_gemma":[0.025254125,0.0003115477,0.0007438794,0.0053070453,0.00043864144,0.0028348956,0.0010244919,0.0012138931,0.0012179287],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014259266,0.0011314289,0.78736377,0.00078687683,0.00047625555,0.00069402345,0.014805056,0.007936661,0.013096523,0.003106773,0.013949834,0.15522689],"study_design_scores_gemma":[0.00006701159,0.000642915,0.6484546,0.00025844324,0.00063135993,0.0017453057,0.015323787,0.29591417,0.012470772,0.008259748,0.016058283,0.00017356238],"about_ca_topic_score_codex":0.004581365,"about_ca_topic_score_gemma":0.008501317,"teacher_disagreement_score":0.00853906,"about_ca_system_score_codex":0.0007007926,"about_ca_system_score_gemma":0.0010414936,"threshold_uncertainty_score":0.017829716},"labels":[],"label_agreement":null},{"id":"W1990080335","doi":"10.4236/jilsa.2011.33015","title":"Insertion of Ontological Knowledge to Improve Automatic Summarization Extraction Methods","year":2011,"lang":"en","type":"article","venue":"Journal of Intelligent Learning Systems and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Abstraction; Set (abstract data type); Information extraction; Function (biology); Artificial intelligence; Natural language processing; Training set; Information retrieval; Machine learning; Selection (genetic algorithm); Data mining; Programming language","score_opus":0.06537092688678746,"score_gpt":0.3557597592021147,"score_spread":0.29038883231532725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1990080335","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042891175,0.00040758034,0.95054704,0.00026481977,0.000053491163,0.00016572108,0.00017880586,0.0035135206,0.0019778665],"genre_scores_gemma":[0.11419594,0.00024043112,0.88285774,0.00007615035,0.000034657674,0.00015151441,0.0009814071,0.00037520568,0.001086902],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99801624,0.0008943458,0.00025367626,0.00030828544,0.00044427282,0.0000832014],"domain_scores_gemma":[0.99510515,0.0025559152,0.00037192352,0.00076701917,0.0011529849,0.00004702462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025784601,0.0010592423,0.00092454016,0.0022975665,0.0007679969,0.0015349612,0.0010500667,0.00075983093,0.0014473503],"category_scores_gemma":[0.012212803,0.00048380665,0.0008130043,0.0017910362,0.00036910054,0.0033935905,0.0012328792,0.0011810169,0.0010237806],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023451104,0.00027356524,0.0017877781,0.00045236092,0.000098653916,0.00013971946,0.0009464424,0.03005795,0.06524741,0.012457723,0.002755496,0.88554853],"study_design_scores_gemma":[0.000121924495,0.00039784558,0.0047759577,0.00013558377,0.00040779755,0.0003453728,0.00087660237,0.7003019,0.22370037,0.02611446,0.04266252,0.0001597147],"about_ca_topic_score_codex":0.0014010986,"about_ca_topic_score_gemma":0.0023283528,"teacher_disagreement_score":0.0025784601,"about_ca_system_score_codex":0.00071169715,"about_ca_system_score_gemma":0.0009990184,"threshold_uncertainty_score":0.01363641},"labels":[],"label_agreement":null},{"id":"W1992795877","doi":"10.3115/1220835.1220896","title":"Language model-based document clustering using random walks","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University; National Science Foundation","keywords":"Cluster analysis; Computer science; Random walk; Document clustering; Representation (politics); Graph; Artificial intelligence; Dimension (graph theory); Hierarchical clustering; Theoretical computer science; Mathematics; Combinatorics; Statistics","score_opus":0.016300133963564334,"score_gpt":0.2558534587018628,"score_spread":0.2395533247382985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1992795877","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004585294,0.00022513767,0.99337834,0.00009354556,0.000032165703,0.000054629938,0.000080843674,0.001136084,0.00041401535],"genre_scores_gemma":[0.14365602,0.00057055976,0.8502968,0.00018973614,0.00013996645,0.00035801998,0.001321993,0.0005875891,0.0028793365],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978404,0.0008473185,0.00011776711,0.0005855763,0.00050005334,0.00010884335],"domain_scores_gemma":[0.99720216,0.0014378174,0.00024029966,0.0005257108,0.00050590874,0.00008817549],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001796068,0.0011367898,0.0018809277,0.0030947176,0.00093069114,0.0020393606,0.0024756088,0.0015549813,0.0015811266],"category_scores_gemma":[0.006878821,0.0006427529,0.0017298541,0.003699772,0.0007992412,0.0035697646,0.0014387066,0.0012918629,0.0020180312],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021147584,0.0002568606,0.0014920157,0.00031570328,0.00028839541,0.00019758145,0.00038801215,0.44255108,0.012648575,0.055792287,0.008877199,0.47698087],"study_design_scores_gemma":[0.000017345288,0.000024607409,0.00010720649,0.000008974756,0.000019005089,0.000057492234,0.000012870762,0.9767671,0.0016459625,0.019832281,0.0014839283,0.000023237071],"about_ca_topic_score_codex":0.0051784893,"about_ca_topic_score_gemma":0.005574237,"teacher_disagreement_score":0.0051784893,"about_ca_system_score_codex":0.00097384804,"about_ca_system_score_gemma":0.0011982722,"threshold_uncertainty_score":0.010296702},"labels":[],"label_agreement":null},{"id":"W1993536130","doi":"10.1145/564376.564438","title":"Using self-supervised word segmentation in Chinese information retrieval","year":2002,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Character (mathematics); Word (group theory); Text segmentation; Natural language processing; Segmentation; Pattern recognition (psychology)","score_opus":0.03339622337105059,"score_gpt":0.2628235046659737,"score_spread":0.22942728129492312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1993536130","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10445216,0.0010780317,0.88565904,0.00019806427,0.000086904234,0.00027147797,0.00014817403,0.004236649,0.003869526],"genre_scores_gemma":[0.55017614,0.0005948682,0.44189715,0.00019856295,0.00021892224,0.00035926382,0.0010607748,0.00034604603,0.005148163],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990408,0.00025205474,0.000100515885,0.0002532542,0.00027613348,0.000077156146],"domain_scores_gemma":[0.9986517,0.00041704992,0.00017174853,0.0002541248,0.00045289643,0.00005251624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073835754,0.0005472685,0.0007344923,0.0020947242,0.0006787533,0.00062451104,0.00081415597,0.00050540693,0.0011927822],"category_scores_gemma":[0.002138242,0.00033807367,0.0006815377,0.001944136,0.000824133,0.0019602925,0.0006299192,0.0004238201,0.0011078882],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025010138,0.00022358901,0.0029375372,0.00036726077,0.000099673,0.0001512361,0.00048566947,0.031407647,0.083860956,0.006181807,0.0053843986,0.86865],"study_design_scores_gemma":[0.00006779557,0.00026255127,0.0036638374,0.000018235743,0.000105143525,0.00026574294,0.0001256125,0.9068989,0.07383682,0.006667986,0.008014505,0.00007291113],"about_ca_topic_score_codex":0.005186132,"about_ca_topic_score_gemma":0.007109409,"teacher_disagreement_score":0.005186132,"about_ca_system_score_codex":0.00056769565,"about_ca_system_score_gemma":0.0012228709,"threshold_uncertainty_score":0.010311902},"labels":[],"label_agreement":null},{"id":"W1995938681","doi":"10.1016/s1364-6613(03)00020-2","title":"Show us the model","year":2003,"lang":"en","type":"review","venue":"Trends in Cognitive Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Scopus; Psychology; Connectionism; Verb; Past tense; Cognitive science; Cognitive psychology; Linguistics; MEDLINE; Neuroscience; Philosophy; Cognition","score_opus":0.44251264546395097,"score_gpt":0.4675551386572874,"score_spread":0.0250424931933364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1995938681","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027110608,0.051245622,0.452864,0.16917285,0.005914418,0.00012654623,0.0032868972,0.0015428652,0.28873616],"genre_scores_gemma":[0.7628149,0.052089557,0.06623878,0.024286898,0.005046398,0.0005855861,0.00269422,0.00036420184,0.08587948],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99957925,0.00015777034,0.000020659098,0.00013136349,0.00007902391,0.000032053937],"domain_scores_gemma":[0.9990337,0.00054374116,0.00008497159,0.00015336984,0.00013878463,0.000045513083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008536458,0.00079111115,0.0006009403,0.0007139971,0.00045507704,0.0022876184,0.00094220333,0.001863432,0.014206777],"category_scores_gemma":[0.0041320072,0.00025705373,0.00054801285,0.00092256634,0.0013436774,0.005556124,0.0008694565,0.0018220657,0.0054372456],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000051019146,0.000024393097,0.0010084384,0.0006233331,0.00008152673,0.00014231497,0.00027371195,0.003121535,0.0007379786,0.8442111,0.07000786,0.079716876],"study_design_scores_gemma":[0.00002717287,0.000023276903,0.0004637271,0.0001476536,0.000074226766,0.00036483066,0.00013840657,0.010880671,0.00034462326,0.8828479,0.10467281,0.000014670749],"about_ca_topic_score_codex":0.0019810796,"about_ca_topic_score_gemma":0.0011309604,"teacher_disagreement_score":0.014206777,"about_ca_system_score_codex":0.00068873336,"about_ca_system_score_gemma":0.0013486583,"threshold_uncertainty_score":0.04752636},"labels":[],"label_agreement":null},{"id":"W1999247825","doi":"10.1016/j.ipm.2009.06.005","title":"Facet-based opinion retrieval from blogs","year":2009,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Facet (psychology); Query expansion; Lexicon; Structuring; Divergence (linguistics); Natural language processing; Artificial intelligence; Linguistics","score_opus":0.01596783699014856,"score_gpt":0.2525663967330099,"score_spread":0.23659855974286134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999247825","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41493392,0.006502839,0.54504055,0.0015424871,0.001036811,0.00074163184,0.010345758,0.0056406492,0.014215335],"genre_scores_gemma":[0.84345555,0.0015757581,0.13142228,0.00024355107,0.00082006236,0.00027038206,0.01538963,0.00033114574,0.0064918003],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99847955,0.0003986364,0.00014336969,0.0002189603,0.0005453012,0.0002142049],"domain_scores_gemma":[0.9955296,0.0021072763,0.0002265513,0.00037807965,0.0015847982,0.0001737657],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016455342,0.0006899859,0.0013852787,0.004607409,0.0009853875,0.002500437,0.0008452063,0.0010205709,0.0029819529],"category_scores_gemma":[0.008520314,0.00027035255,0.0014551482,0.0038686208,0.0003712727,0.002469224,0.0011330578,0.00078659714,0.0027135566],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031810019,0.0006652435,0.024310894,0.0016708174,0.0007031788,0.0008889903,0.0018655453,0.015895635,0.10257599,0.006385695,0.061744306,0.7801127],"study_design_scores_gemma":[0.00021759661,0.0007455901,0.0334668,0.00019423559,0.0008390757,0.0010519141,0.0019798903,0.87589777,0.037939094,0.025311612,0.022173312,0.00018316746],"about_ca_topic_score_codex":0.0040575988,"about_ca_topic_score_gemma":0.008556348,"teacher_disagreement_score":0.004607409,"about_ca_system_score_codex":0.00063221855,"about_ca_system_score_gemma":0.0009198472,"threshold_uncertainty_score":0.009975672},"labels":[],"label_agreement":null},{"id":"W1999484828","doi":"10.1109/icsm.2013.72","title":"On the Personality Traits of StackOverflow Users","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":111,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Reputation; Voting; Big Five personality traits; Personality; Computer science; World Wide Web; Psychology; Social psychology; Sociology; Political science; Social science","score_opus":0.03420112286434147,"score_gpt":0.2327080531954025,"score_spread":0.19850693033106104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999484828","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9992828,0.000032606687,0.00008632085,0.000043711912,0.000003519961,0.0000041925773,0.00003284698,0.0000029656762,0.0005110163],"genre_scores_gemma":[0.99954647,0.0000287002,0.00007737497,0.000017173632,0.0000056456565,0.000003904882,0.00003827998,0.0000018282238,0.00028064614],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.999466,0.00017287016,0.000048044338,0.00006260435,0.00016618927,0.00008426465],"domain_scores_gemma":[0.9905758,0.00325863,0.003248993,0.00040850803,0.0011193996,0.0013886306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00082877267,0.00019901645,0.00019196869,0.0012344186,0.0005016252,0.0011561458,0.00013749709,0.00028569682,0.001537639],"category_scores_gemma":[0.009783252,0.00014744682,0.000234022,0.0005891959,0.0003347327,0.0006372595,0.00052362395,0.00041820682,0.00040041027],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000112562135,0.000062336774,0.98494786,0.000015512811,0.000028442068,0.00012581416,0.005183906,0.000065450295,0.0010544326,0.00004558734,0.00021264463,0.008145521],"study_design_scores_gemma":[0.0000031856323,0.00009579444,0.9932448,0.000009518299,0.00001321592,0.00023230449,0.0049118167,0.000712094,0.00019209302,0.00008378719,0.00048706416,0.000014212783],"about_ca_topic_score_codex":0.0019456224,"about_ca_topic_score_gemma":0.0022768292,"teacher_disagreement_score":0.0019456224,"about_ca_system_score_codex":0.00020734445,"about_ca_system_score_gemma":0.0001229546,"threshold_uncertainty_score":0.0051439404},"labels":[],"label_agreement":null},{"id":"W2000231926","doi":"10.1109/cogsima.2014.6816558","title":"Textual risk mining for maritime situational awareness","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Larus Technologies (Canada); University of Ottawa","funders":"National Geospatial-Intelligence Agency","keywords":"Situation awareness; Computer science; Situational ethics; Data science; Engineering; Psychology","score_opus":0.02708779864761301,"score_gpt":0.26431921952053467,"score_spread":0.23723142087292165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000231926","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040585667,0.0009842814,0.9209952,0.0009422466,0.00013422781,0.00048635353,0.010836259,0.019031027,0.0060047912],"genre_scores_gemma":[0.38185048,0.00065964815,0.5873253,0.00031140068,0.00026834666,0.00052711566,0.02464894,0.0009178103,0.0034910415],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99775994,0.00058139476,0.00033874478,0.000570667,0.0006770546,0.00007222246],"domain_scores_gemma":[0.9946629,0.0029421838,0.0008249705,0.0007042417,0.00072720746,0.00013856117],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001283803,0.0011006088,0.0006260334,0.006345887,0.000640083,0.0018773789,0.001149869,0.0007462747,0.0045951894],"category_scores_gemma":[0.009187485,0.0004220442,0.000987203,0.0032145493,0.00045824936,0.0036159772,0.0016659953,0.000868067,0.0025450422],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006445919,0.00052579015,0.01337111,0.0021134792,0.0002273882,0.0015347386,0.0017622205,0.028092956,0.039863937,0.028295962,0.024528896,0.85903895],"study_design_scores_gemma":[0.000098281416,0.00023478032,0.012873013,0.0004079073,0.00033347905,0.0018709959,0.001347273,0.7425241,0.04908402,0.09745315,0.093608335,0.00016469309],"about_ca_topic_score_codex":0.0017485081,"about_ca_topic_score_gemma":0.0020379955,"teacher_disagreement_score":0.006345887,"about_ca_system_score_codex":0.0007219979,"about_ca_system_score_gemma":0.0010610269,"threshold_uncertainty_score":0.015372455},"labels":[],"label_agreement":null},{"id":"W2000545282","doi":"10.1145/383952.384024","title":"Exploiting redundancy in question answering","year":2001,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":237,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Question answering; Computer science; Redundancy (engineering); Information retrieval; Questions and answers; Artificial intelligence","score_opus":0.029088710983966825,"score_gpt":0.26895172356739283,"score_spread":0.23986301258342602,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000545282","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09009948,0.005122905,0.88830286,0.0024204822,0.0002111034,0.00056101993,0.0014278659,0.006452028,0.005402248],"genre_scores_gemma":[0.5268651,0.0016144556,0.46080953,0.00087626546,0.00081609807,0.000638994,0.0049517523,0.0005373049,0.0028904204],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98390573,0.009558165,0.0009483606,0.002059374,0.0027759902,0.00075226754],"domain_scores_gemma":[0.9560141,0.034657694,0.0017102145,0.004651231,0.002615417,0.00035138224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009978613,0.0016219467,0.002615532,0.008480045,0.0018325791,0.0025211088,0.0027670667,0.00243824,0.003470486],"category_scores_gemma":[0.053000677,0.0010203978,0.0018764854,0.0053149145,0.0016153335,0.007497381,0.004567026,0.001738486,0.0019335587],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012480464,0.0005393752,0.013169685,0.002280975,0.00045383963,0.0010757694,0.00551295,0.031039156,0.02550307,0.041661344,0.027056715,0.85045904],"study_design_scores_gemma":[0.00021767346,0.0005838433,0.0077863107,0.00031640116,0.00059684034,0.0024212294,0.0015986419,0.6524826,0.028185384,0.26044917,0.045165684,0.00019618217],"about_ca_topic_score_codex":0.002561537,"about_ca_topic_score_gemma":0.0022234365,"teacher_disagreement_score":0.009978613,"about_ca_system_score_codex":0.0010055845,"about_ca_system_score_gemma":0.00119321,"threshold_uncertainty_score":0.052772522},"labels":[],"label_agreement":null},{"id":"W2001214180","doi":"10.5539/cis.v3n1p168","title":"Automatic Recognition of Focus and Interrogative Word in Chinese Question for Classification","year":2010,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Interrogative; Computer science; Focus (optics); Artificial intelligence; Natural language processing; Interrogative word; CRFS; Conditional random field; Word (group theory); Parsing; Part of speech; Text segmentation; Segmentation; Linguistics","score_opus":0.023976777368841653,"score_gpt":0.2882214725990372,"score_spread":0.26424469523019556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2001214180","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75648934,0.0010522598,0.22329569,0.00054653396,0.00017259544,0.0006250507,0.003682354,0.0063362396,0.0077998713],"genre_scores_gemma":[0.92697865,0.00020128419,0.06309111,0.00011332562,0.00006136555,0.00025123789,0.006205612,0.0001496138,0.0029477878],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987852,0.00030806582,0.00011945629,0.00041037262,0.00019030205,0.000186706],"domain_scores_gemma":[0.995261,0.0017728513,0.00037916476,0.00051163253,0.0018691038,0.00020620297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017918258,0.00067534164,0.00074583036,0.0025793298,0.00065975514,0.0007451657,0.00073370535,0.0006421764,0.0024175213],"category_scores_gemma":[0.004786194,0.00022930549,0.0006248075,0.0011174962,0.00048503536,0.0020179516,0.0009708784,0.00053686736,0.0012058951],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087255734,0.00029746996,0.080245055,0.0009481461,0.00008844662,0.00071942987,0.006121106,0.0039489367,0.24715914,0.0074037868,0.020709718,0.6314863],"study_design_scores_gemma":[0.00013968023,0.00067182846,0.26890752,0.00011509565,0.00034724342,0.0016348404,0.0033879848,0.4300582,0.25377765,0.008299848,0.03241482,0.00024532067],"about_ca_topic_score_codex":0.009937418,"about_ca_topic_score_gemma":0.005305432,"teacher_disagreement_score":0.009937418,"about_ca_system_score_codex":0.00069016777,"about_ca_system_score_gemma":0.0010454302,"threshold_uncertainty_score":0.019759119},"labels":[],"label_agreement":null},{"id":"W2002358989","doi":"10.1080/17470210903267417","title":"Applying an exemplar model to the artificial-grammar task: String completion and performance on individual items","year":2009,"lang":"en","type":"article","venue":"Quarterly Journal of Experimental Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Manitoba","funders":"","keywords":"Grammar; String (physics); Task (project management); Computer science; Natural language processing; Artificial intelligence; Psychology; Linguistics; Mathematics","score_opus":0.07143995998942312,"score_gpt":0.3442594578043688,"score_spread":0.27281949781494563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002358989","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99363285,0.00004400514,0.004935449,0.00010053261,0.00002346247,0.00006727834,0.00018868192,0.00006823733,0.0009395784],"genre_scores_gemma":[0.9918833,0.000042698455,0.005670773,0.00010818644,0.000025117291,0.00014395523,0.0006386845,0.000062008505,0.0014252928],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980001,0.0008432651,0.0001545924,0.00052618166,0.00038888442,0.00008695351],"domain_scores_gemma":[0.9564461,0.028076237,0.0044923606,0.008218592,0.0011888165,0.0015778868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006755332,0.0007440733,0.0010111522,0.0006548086,0.0003279846,0.0015263952,0.0018249352,0.0021465172,0.0034535055],"category_scores_gemma":[0.03762285,0.0006090193,0.00054174947,0.00056363473,0.0009605862,0.003017126,0.0015134974,0.0019347933,0.0009248931],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.030031422,0.032573562,0.57165736,0.00067472825,0.0021042766,0.0010576705,0.010468876,0.03745543,0.15715434,0.011715456,0.0074219275,0.13768493],"study_design_scores_gemma":[0.0011689997,0.01308544,0.6550245,0.000061104205,0.00044282828,0.0012431549,0.00067009917,0.26905176,0.022207607,0.03443776,0.0022953644,0.00031129387],"about_ca_topic_score_codex":0.0021532767,"about_ca_topic_score_gemma":0.0014072855,"teacher_disagreement_score":0.006755332,"about_ca_system_score_codex":0.00043339844,"about_ca_system_score_gemma":0.0004357155,"threshold_uncertainty_score":0.03572607},"labels":[],"label_agreement":null},{"id":"W200414342","doi":"10.1007/978-3-642-24571-8_2","title":"A Pattern-Based Model for Generating Text to Express Emotion","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Sentence; Representation (politics); Natural language processing; Artificial intelligence; Class (philosophy)","score_opus":0.04712581315086143,"score_gpt":0.25549651621048153,"score_spread":0.2083707030596201,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W200414342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010986312,0.00013311945,0.98104596,0.00036818825,0.00010929868,0.00027978863,0.0013595464,0.0033836025,0.0023342283],"genre_scores_gemma":[0.2582725,0.00033005793,0.72537905,0.00025540788,0.00011354017,0.0010558142,0.0035131036,0.00056024455,0.010520376],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99929225,0.000180128,0.00006596004,0.0002472333,0.00017179314,0.000042529125],"domain_scores_gemma":[0.99832815,0.0009119519,0.00009561645,0.00016769099,0.0004404044,0.000056158195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009872302,0.0007281139,0.000479795,0.0010835242,0.00039549617,0.0012628129,0.0014405504,0.0010684661,0.0058473293],"category_scores_gemma":[0.004564523,0.00037261247,0.0010078959,0.00095515704,0.0003844046,0.002099861,0.0006053656,0.0009927687,0.0030226286],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010869431,0.00041465374,0.006206022,0.00071291864,0.00021698793,0.000826814,0.001302905,0.077818766,0.058386322,0.04927116,0.03486234,0.7688942],"study_design_scores_gemma":[0.000053036198,0.0001271211,0.0011243221,0.000035498455,0.00009321627,0.0002863208,0.00009765405,0.94688934,0.011656953,0.029710528,0.009894577,0.00003139673],"about_ca_topic_score_codex":0.003753086,"about_ca_topic_score_gemma":0.0042014,"teacher_disagreement_score":0.0058473293,"about_ca_system_score_codex":0.00064887304,"about_ca_system_score_gemma":0.00065638975,"threshold_uncertainty_score":0.01956129},"labels":[],"label_agreement":null},{"id":"W2006258148","doi":"10.1075/term.14.1.02aug","title":"Pattern-based approaches to semantic relation extraction","year":2008,"lang":"en","type":"article","venue":"Terminology International Journal of Theoretical and Applied Issues in Specialized Communication","topic":"Topic Modeling","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Université de Nantes; Université de Montréal; Université de Lyon; Centre National de la Recherche Scientifique; University of Ottawa","keywords":"Relation (database); Computer science; Relationship extraction; Semantic relation; Natural language processing; Artificial intelligence; Information retrieval; Data mining; Psychology","score_opus":0.06600968974962766,"score_gpt":0.3079913575456,"score_spread":0.2419816677959723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006258148","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014875736,0.00238405,0.94109464,0.005351533,0.0014552941,0.0010518535,0.005663806,0.008282216,0.019840937],"genre_scores_gemma":[0.12289606,0.0026156968,0.8275974,0.0006628517,0.00078136765,0.0007811877,0.014778083,0.002524661,0.027362764],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98941576,0.002851561,0.001542936,0.001391241,0.0044916407,0.0003068396],"domain_scores_gemma":[0.96482444,0.013562396,0.0017065635,0.005708735,0.0135845635,0.0006132495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007754899,0.0009305692,0.0015298116,0.0107106,0.0026159426,0.006937456,0.0028426298,0.002361653,0.026619358],"category_scores_gemma":[0.046932817,0.0012357302,0.0019155907,0.011591499,0.0013296244,0.010189011,0.0031337661,0.0022796884,0.016629765],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031295986,0.00019165201,0.0041950108,0.0019252776,0.00023303648,0.00080532685,0.0013372694,0.0022835738,0.017801031,0.045827564,0.08950471,0.83558273],"study_design_scores_gemma":[0.00019380795,0.00023642005,0.008706707,0.00097031594,0.0005679874,0.0030792041,0.0028194978,0.106247984,0.049425878,0.23832329,0.58917534,0.00025354675],"about_ca_topic_score_codex":0.0016527175,"about_ca_topic_score_gemma":0.0032393562,"teacher_disagreement_score":0.026619358,"about_ca_system_score_codex":0.000989973,"about_ca_system_score_gemma":0.0034455261,"threshold_uncertainty_score":0.08905059},"labels":[],"label_agreement":null},{"id":"W2006329944","doi":"10.1037/a0027023","title":"An exemplar model of performance in the artificial grammar task: Holographic representation.","year":2012,"lang":"en","type":"article","venue":"Canadian Journal of Experimental Psychology/Revue canadienne de psychologie expérimentale","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Grammar; Memory model; Linguistics","score_opus":0.08623202579178714,"score_gpt":0.3302681238390607,"score_spread":0.24403609804727355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2006329944","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18823084,0.00078583084,0.78318965,0.0028667573,0.00014962559,0.0003061656,0.0011631717,0.0009099894,0.022397982],"genre_scores_gemma":[0.90246075,0.0003608732,0.08855262,0.0002723236,0.000060211045,0.00042613244,0.0006997922,0.00013593481,0.0070314007],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99906033,0.0003588885,0.000050098588,0.0002799244,0.00015677093,0.000093936644],"domain_scores_gemma":[0.9938446,0.0037976592,0.00064115826,0.0010198778,0.0004466997,0.00025004864],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024226902,0.0005273897,0.0008129478,0.0011301995,0.00038174106,0.0019377032,0.0030830076,0.0017502209,0.006811082],"category_scores_gemma":[0.013057204,0.0004528603,0.0016785053,0.00097039476,0.0013575397,0.0044985977,0.0013019771,0.0013653316,0.0010947854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000624122,0.00060479034,0.019221013,0.0003647147,0.00043509327,0.0006410463,0.002615576,0.14620963,0.010144776,0.7229101,0.005603984,0.09062508],"study_design_scores_gemma":[0.00007116187,0.0002189508,0.006602067,0.000020806681,0.00007935319,0.00052114774,0.00013275555,0.6709925,0.0011665305,0.31706414,0.0030792735,0.000051349474],"about_ca_topic_score_codex":0.004436098,"about_ca_topic_score_gemma":0.0020291333,"teacher_disagreement_score":0.006811082,"about_ca_system_score_codex":0.0011101104,"about_ca_system_score_gemma":0.0010751141,"threshold_uncertainty_score":0.022785366},"labels":[],"label_agreement":null},{"id":"W2008283248","doi":"10.1007/s10791-013-9220-9","title":"Latent word context model for information retrieval","year":2013,"lang":"en","type":"article","venue":"Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Ministry of Education, Culture, Sports, Science and Technology","keywords":"Computer science; Latent Dirichlet allocation; Word (group theory); Context (archaeology); Information retrieval; Natural language processing; Latent semantic analysis; Search engine indexing; Topic model; Relevance (law); Artificial intelligence; Probabilistic latent semantic analysis; Linguistics","score_opus":0.02483661079381079,"score_gpt":0.23769802741019738,"score_spread":0.21286141661638658,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008283248","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010183527,0.0063549,0.97846425,0.00095682626,0.00020232525,0.00008465172,0.0009955954,0.0011051462,0.0016527383],"genre_scores_gemma":[0.5755696,0.009994214,0.38523626,0.0007698655,0.0012872367,0.0010599986,0.006459353,0.0006990386,0.01892441],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99740404,0.0014005876,0.00017201081,0.0004671043,0.00037473213,0.00018140729],"domain_scores_gemma":[0.99616915,0.0026517718,0.00023809951,0.00047034718,0.0003890096,0.00008161241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028986835,0.0010552829,0.0022142555,0.0025204218,0.0007558711,0.002495525,0.0024519907,0.0022190518,0.004332417],"category_scores_gemma":[0.011401789,0.00077275647,0.0016735547,0.0043547694,0.0009203355,0.0058719846,0.0012929911,0.0029379912,0.0031407601],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008004137,0.00050865,0.0032404286,0.0012709951,0.0006560082,0.00037715433,0.00078132545,0.2073346,0.0076058484,0.39307812,0.02753184,0.3568146],"study_design_scores_gemma":[0.00004969746,0.00005722413,0.0006815299,0.00005740674,0.00013578428,0.00010273981,0.00005152856,0.80075395,0.0008455367,0.19104514,0.006172834,0.00004664924],"about_ca_topic_score_codex":0.0075179352,"about_ca_topic_score_gemma":0.0075265607,"teacher_disagreement_score":0.0075179352,"about_ca_system_score_codex":0.0014092543,"about_ca_system_score_gemma":0.0016894806,"threshold_uncertainty_score":0.015329838},"labels":[],"label_agreement":null},{"id":"W2010181950","doi":"10.1145/1277741.1277934","title":"Power and bias of subset pooling strategies","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Pooling; Statistics; Computer science; Relevance (law); Rank (graph theory); Adjudication; Statistical power; Power (physics); Yield (engineering); Correlation; Artificial intelligence; Data mining; Econometrics; Mathematics; Combinatorics; Geometry","score_opus":0.0381045557189526,"score_gpt":0.2734510117031871,"score_spread":0.23534645598423448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2010181950","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025290487,0.0013434098,0.9681892,0.0008796989,0.000097026685,0.0004738452,0.00024439354,0.00052795873,0.0029540334],"genre_scores_gemma":[0.57047814,0.001317581,0.4203856,0.00078620773,0.0006785502,0.002256503,0.0007396307,0.000484324,0.0028734996],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9342548,0.04934733,0.0028492396,0.0051257345,0.007369982,0.0010528442],"domain_scores_gemma":[0.71044695,0.25272045,0.0065560755,0.021253902,0.008185081,0.0008375549],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09740604,0.0027927135,0.004384095,0.003996516,0.0012945522,0.0045419727,0.003269843,0.002898383,0.0031794612],"category_scores_gemma":[0.32258084,0.0013912215,0.0023403305,0.0036134818,0.0032797763,0.007956355,0.0062008104,0.0022743498,0.0010443153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020509635,0.00023392834,0.017444642,0.0012509489,0.0031278345,0.00030786535,0.0014370452,0.14037277,0.007429585,0.1267813,0.0067354916,0.69282764],"study_design_scores_gemma":[0.0004336611,0.0010628707,0.006767053,0.00034265005,0.0012462575,0.0005941857,0.00038284296,0.6255644,0.014817605,0.34234956,0.006290401,0.00014852922],"about_ca_topic_score_codex":0.0018175219,"about_ca_topic_score_gemma":0.0012234126,"teacher_disagreement_score":0.09740604,"about_ca_system_score_codex":0.0012938849,"about_ca_system_score_gemma":0.0030882808,"threshold_uncertainty_score":0.5151385},"labels":[],"label_agreement":null},{"id":"W2010612305","doi":"10.1145/1621995.1622025","title":"Rhetorical models for computational systems","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Rhetorical question; Computer science; Scholarship; Natural language processing; Artificial intelligence; Documentation; Field (mathematics); Linguistics; Programming language; Mathematics; Political science","score_opus":0.0697503413865748,"score_gpt":0.288834156603689,"score_spread":0.2190838152171142,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2010612305","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0067506344,0.0048027444,0.9341169,0.008617313,0.0003139332,0.000116049654,0.0006154881,0.00080822676,0.043858726],"genre_scores_gemma":[0.46283376,0.005815851,0.5035625,0.0015358541,0.0013266755,0.0011380045,0.001594853,0.00054505566,0.021647414],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99784243,0.001207597,0.0001377867,0.00033162648,0.00037810463,0.00010241341],"domain_scores_gemma":[0.99207217,0.006188998,0.0002946302,0.00086722913,0.00040507986,0.00017191857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036117414,0.0009979878,0.00096355146,0.0022364054,0.0015167651,0.0052477606,0.0019210962,0.0022199356,0.010756431],"category_scores_gemma":[0.01167293,0.00056994485,0.0016534786,0.0017407599,0.0044877836,0.008231417,0.0026684261,0.0032608511,0.0023405442],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000004448112,0.000007639982,0.0000792312,0.00004829446,0.000008782426,0.000019607112,0.00012997318,0.008913778,0.00005615189,0.98440695,0.0014031255,0.004922046],"study_design_scores_gemma":[0.000004436248,0.0000030028295,0.000021121128,0.000015998581,0.000002899043,0.000011064142,0.000026908117,0.04039316,0.000037254013,0.9520672,0.007412945,0.000004088762],"about_ca_topic_score_codex":0.0027718025,"about_ca_topic_score_gemma":0.002615991,"teacher_disagreement_score":0.010756431,"about_ca_system_score_codex":0.0027840992,"about_ca_system_score_gemma":0.0017194337,"threshold_uncertainty_score":0.03598386},"labels":[],"label_agreement":null},{"id":"W2013355692","doi":"10.1016/j.jbi.2012.11.006","title":"Detecting concept relations in clinical text: Insights from a state-of-the-art model","year":2013,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"U.S. National Library of Medicine","keywords":"Computer science; Variety (cybernetics); Construct (python library); Artificial intelligence; Feature (linguistics); Natural language processing; Relationship extraction; Feature engineering; Information retrieval; State (computer science); Kernel (algebra); Data science; Information extraction; Machine learning; Deep learning; Linguistics; Programming language; Mathematics","score_opus":0.02697396741177177,"score_gpt":0.2853127824758678,"score_spread":0.258338815064096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2013355692","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16267882,0.01391221,0.8080517,0.0051349807,0.00035847788,0.00049509684,0.0042353473,0.0022200702,0.0029132862],"genre_scores_gemma":[0.76329833,0.0047159814,0.22061038,0.0007389861,0.0009705479,0.0005004598,0.00691611,0.00033369777,0.001915537],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9929577,0.0024982728,0.0010486572,0.0019835888,0.0011970486,0.00031476683],"domain_scores_gemma":[0.95451826,0.039837115,0.0012871883,0.0018067808,0.0021013422,0.00044942944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011385841,0.0014146368,0.0023403862,0.0069302456,0.0013490424,0.0068638884,0.0029974454,0.002794406,0.0018837668],"category_scores_gemma":[0.033821147,0.0007706937,0.003194643,0.0039598,0.0014220468,0.0077575115,0.0022713772,0.0028206704,0.0014354925],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026075672,0.0012758706,0.06704476,0.003459552,0.0019574908,0.0010939483,0.0043241456,0.12399238,0.0195759,0.03795977,0.018013744,0.7186948],"study_design_scores_gemma":[0.00007166453,0.00020846428,0.008572708,0.00019478746,0.0005934847,0.0005947663,0.00030474807,0.93407136,0.002858664,0.046800215,0.0056565935,0.000072525116],"about_ca_topic_score_codex":0.0060514812,"about_ca_topic_score_gemma":0.0070418622,"teacher_disagreement_score":0.011385841,"about_ca_system_score_codex":0.0012339754,"about_ca_system_score_gemma":0.0025000575,"threshold_uncertainty_score":0.060214818},"labels":[],"label_agreement":null},{"id":"W2013809972","doi":"10.1109/icosc.2015.7050809","title":"Aligning automatically generated questions to instructor goals and learner behaviour","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Ranking (information retrieval); Process (computing); Quality (philosophy); Artificial intelligence; Natural language processing; Information retrieval","score_opus":0.042032487168416746,"score_gpt":0.2841595463794267,"score_spread":0.24212705921100997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2013809972","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16261417,0.00037363163,0.81584966,0.00032188636,0.00012164773,0.0016876147,0.0014249043,0.015095263,0.0025111572],"genre_scores_gemma":[0.4711737,0.00018034982,0.51562405,0.00016442144,0.00011333784,0.0013759625,0.007928931,0.0011874583,0.0022517461],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9849577,0.010087211,0.00089110323,0.0017703739,0.0020398041,0.000253884],"domain_scores_gemma":[0.9174507,0.058284756,0.004681993,0.006290596,0.0125723705,0.0007196247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010828675,0.0013049783,0.0009784289,0.002425649,0.00045470844,0.0019984161,0.0014818395,0.0016150321,0.0021256164],"category_scores_gemma":[0.07050514,0.00041404113,0.00079564465,0.0012405083,0.00055150193,0.0020686,0.0013356229,0.001161312,0.0013825439],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014700103,0.0017105509,0.037996247,0.0018095382,0.00033609165,0.00033650122,0.0051702117,0.037904464,0.13728686,0.0044663195,0.00906073,0.76245254],"study_design_scores_gemma":[0.00032759024,0.0016610518,0.050604373,0.00018661526,0.00026099544,0.0005492965,0.0013094866,0.7363142,0.17521125,0.0120977145,0.021209896,0.0002675596],"about_ca_topic_score_codex":0.0016696849,"about_ca_topic_score_gemma":0.001937695,"teacher_disagreement_score":0.010828675,"about_ca_system_score_codex":0.0007805965,"about_ca_system_score_gemma":0.0010785423,"threshold_uncertainty_score":0.057268202},"labels":[],"label_agreement":null},{"id":"W2015722806","doi":"10.1007/s10503-008-9104-0","title":"Evaluating Corroborative Evidence","year":2008,"lang":"en","type":"article","venue":"Argumentation","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Political communication; Communication studies; Political science; Psychology; Computer science; Sociology; Social science; Politics; Law","score_opus":0.23986453346287462,"score_gpt":0.4065833421626628,"score_spread":0.16671880869978817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2015722806","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28686145,0.011588611,0.60876536,0.024012353,0.0013124861,0.00131025,0.002181191,0.0014811793,0.062487062],"genre_scores_gemma":[0.84119046,0.0013936164,0.15128203,0.0005088197,0.00077534036,0.00035790927,0.0014889438,0.000150254,0.002852575],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.89629227,0.06456127,0.0066609886,0.006355863,0.024717076,0.0014125331],"domain_scores_gemma":[0.35558704,0.5628299,0.025153657,0.03107917,0.022603367,0.0027467648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.061319016,0.0020378684,0.0021758967,0.014513301,0.0032077192,0.011635225,0.0040621837,0.007995498,0.011555404],"category_scores_gemma":[0.47731322,0.0015115228,0.0021934884,0.0061355005,0.004431642,0.016601855,0.0077173356,0.0046247547,0.0016985566],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003873946,0.0011509834,0.048943847,0.0045986683,0.0023685824,0.0024232746,0.0097265085,0.028822163,0.0065757823,0.31841236,0.017620228,0.5554835],"study_design_scores_gemma":[0.00053395174,0.0008130385,0.012231815,0.0025282516,0.0017326933,0.0014489405,0.004370813,0.2133702,0.016452165,0.69912976,0.047165237,0.00022321404],"about_ca_topic_score_codex":0.0012641498,"about_ca_topic_score_gemma":0.00196876,"teacher_disagreement_score":0.061319016,"about_ca_system_score_codex":0.0027338215,"about_ca_system_score_gemma":0.0042321766,"threshold_uncertainty_score":0.3242898},"labels":[],"label_agreement":null},{"id":"W2016568784","doi":"10.3115/1118958.1118968","title":"Answering clinical questions with role identification","year":2003,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; University of Toronto","keywords":"Question answering; Computer science; Identification (biology); Context (archaeology); Domain (mathematical analysis); Natural language; Information retrieval; Data science; Natural language processing; Artificial intelligence","score_opus":0.032749389575024886,"score_gpt":0.30886842886973054,"score_spread":0.27611903929470566,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2016568784","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028031925,0.0021416922,0.9507628,0.006159797,0.00025597741,0.00079649925,0.0013437921,0.002497954,0.008009565],"genre_scores_gemma":[0.28111094,0.0012170691,0.70683444,0.0015014678,0.0006557238,0.0007973217,0.0034853325,0.00018109231,0.0042166333],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9917965,0.0054604462,0.00037108798,0.0012132706,0.00083382655,0.0003249002],"domain_scores_gemma":[0.9809428,0.014190871,0.0010730582,0.0017452879,0.0015086933,0.0005393162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0085602,0.0010238644,0.0008709706,0.0033893774,0.0009746036,0.0031930245,0.0017718424,0.0018664304,0.0065836883],"category_scores_gemma":[0.025941534,0.00039937734,0.0015341626,0.0014646839,0.0012206045,0.0048369807,0.0025121146,0.00178309,0.0028583247],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010284625,0.0010964028,0.026860794,0.0015569909,0.00021396817,0.0013390431,0.005518796,0.015500698,0.024688682,0.082916856,0.03327237,0.8060069],"study_design_scores_gemma":[0.00032154273,0.00066231156,0.010830002,0.0004317284,0.0004111198,0.004138828,0.0039830552,0.33843213,0.035976548,0.4651502,0.13947567,0.00018692226],"about_ca_topic_score_codex":0.0010830557,"about_ca_topic_score_gemma":0.0016354067,"teacher_disagreement_score":0.0085602,"about_ca_system_score_codex":0.0007899277,"about_ca_system_score_gemma":0.0012715994,"threshold_uncertainty_score":0.04527116},"labels":[],"label_agreement":null},{"id":"W2018189188","doi":"10.1073/pnas.1219674110","title":"Limits in decision making arise from limits in memory retrieval","year":2013,"lang":"en","type":"article","venue":"Proceedings of the National Academy of Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Institute of Mental Health","keywords":"Computer science; Idealization; Probabilistic logic; Process (computing); Machine learning; Set (abstract data type); Artificial intelligence; Noise (video); Rationality; Contrast (vision); Sample (material); Test (biology)","score_opus":0.059567304085905264,"score_gpt":0.3109434555717325,"score_spread":0.2513761514858272,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2018189188","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58487225,0.003031558,0.35711902,0.008057047,0.000100916805,0.000116656534,0.00037230877,0.0005618468,0.04576847],"genre_scores_gemma":[0.9666557,0.0005313539,0.030282337,0.0006723563,0.00006562211,0.00013588183,0.0001677267,0.000092293994,0.0013967431],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.99037206,0.0028938465,0.0008916373,0.0025526283,0.0025465763,0.0007432983],"domain_scores_gemma":[0.92690676,0.052955143,0.005727692,0.010194738,0.0026850733,0.0015306121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009320477,0.00070978404,0.0012985868,0.0012506716,0.00078772934,0.0064922725,0.0022022924,0.0022668561,0.0029521429],"category_scores_gemma":[0.09583497,0.0011263518,0.0010145429,0.00097065925,0.007832335,0.007958222,0.0036501412,0.0034449724,0.0009215008],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010542268,0.0004947432,0.047224913,0.00064842607,0.0007012803,0.0011017263,0.0061919554,0.082315795,0.021201044,0.6223238,0.003908862,0.21283323],"study_design_scores_gemma":[0.00006859194,0.00013412932,0.0139488075,0.00007388192,0.00005736583,0.00028706723,0.0006142572,0.051163483,0.0033806209,0.92756945,0.0026093025,0.000093021714],"about_ca_topic_score_codex":0.0023145382,"about_ca_topic_score_gemma":0.0010963827,"teacher_disagreement_score":0.009320477,"about_ca_system_score_codex":0.0018704107,"about_ca_system_score_gemma":0.0011363784,"threshold_uncertainty_score":0.049292028},"labels":[],"label_agreement":null},{"id":"W2020278455","doi":"10.1075/li.30.1.03nad","title":"A survey of named entity recognition and classification","year":2007,"lang":"en","type":"article","venue":"Lingvisticae Investigationes","topic":"Topic Modeling","field":"Computer Science","cited_by":2488,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Natural Environment Research Council; Alfred P. Sloan Foundation","keywords":"Computer science; Natural language processing; Matching (statistics); Artificial intelligence; Field (mathematics); Information retrieval; Named entity; Key (lock); Word (group theory); Linguistics; Mathematics","score_opus":0.1391236913969597,"score_gpt":0.291564756940622,"score_spread":0.1524410655436623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2020278455","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006312263,0.60615045,0.33282706,0.0078098285,0.0026592254,0.0005552437,0.0063029574,0.0056854216,0.031697642],"genre_scores_gemma":[0.03708992,0.679804,0.23348543,0.0033968256,0.0040064123,0.00057031866,0.024376426,0.0008212925,0.016449358],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99275136,0.0017782812,0.000977919,0.001259781,0.0029213314,0.00031137845],"domain_scores_gemma":[0.983254,0.008957142,0.00081395404,0.0018879715,0.0047197067,0.00036729863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007817268,0.0018341141,0.0030622433,0.0143100815,0.0014219526,0.004739904,0.0039729625,0.0020924648,0.008388628],"category_scores_gemma":[0.01780604,0.0008817977,0.0018035122,0.027164016,0.0008828376,0.017031055,0.0021527226,0.0019936904,0.009767558],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000054540727,0.0000934599,0.0025922093,0.0029134504,0.000091422015,0.00011069061,0.0001230192,0.0016026766,0.0006050287,0.009097715,0.050344817,0.9323709],"study_design_scores_gemma":[0.000015284992,0.00012942332,0.0045723035,0.0018127472,0.00015589874,0.0012744152,0.0005029165,0.017446829,0.0034089417,0.01920094,0.95136917,0.000111291556],"about_ca_topic_score_codex":0.003658784,"about_ca_topic_score_gemma":0.0028822718,"teacher_disagreement_score":0.0143100815,"about_ca_system_score_codex":0.0015378263,"about_ca_system_score_gemma":0.0030041588,"threshold_uncertainty_score":0.04134214},"labels":[],"label_agreement":null},{"id":"W2025486736","doi":"10.1177/154193120204602412","title":"Evaluating the Content and Usability of an Experimental Text Summarization System and Three Web-Based Search Engines","year":2002,"lang":"en","type":"article","venue":"Proceedings of the Human Factors and Ergonomics Society Annual Meeting","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Automatic summarization; Computer science; Think aloud protocol; Information retrieval; Usability; Search engine; World Wide Web; Interface (matter); Subject (documents); Human–computer interaction","score_opus":0.07916089657933037,"score_gpt":0.28262016902202686,"score_spread":0.2034592724426965,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025486736","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9940117,0.00020777702,0.0037514819,0.000044270273,0.000022431424,0.00088552537,0.00016973865,0.00028027475,0.00062688324],"genre_scores_gemma":[0.9538013,0.0004073782,0.039609518,0.00015437415,0.00011156329,0.001887707,0.0016815204,0.00017024387,0.0021765118],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9933106,0.0039534094,0.0010908621,0.00064323284,0.0007596618,0.00024236963],"domain_scores_gemma":[0.92488855,0.06405516,0.0020417036,0.0016514165,0.0061746314,0.0011885483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008865261,0.0011320577,0.0012609864,0.0011780906,0.00077666185,0.0013349087,0.0010668169,0.0012593811,0.0019504552],"category_scores_gemma":[0.043355476,0.0005973618,0.0006717901,0.0010987996,0.00067423226,0.0019551914,0.0010592362,0.00049161125,0.00052875275],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.048523493,0.03512664,0.06590678,0.009712705,0.0013863166,0.0015251912,0.03258881,0.012571551,0.35997614,0.0007459623,0.0043481984,0.4275882],"study_design_scores_gemma":[0.011903656,0.22250283,0.343876,0.00051055837,0.00518675,0.0016698773,0.011291525,0.10790164,0.27840608,0.00093657075,0.015070843,0.0007437016],"about_ca_topic_score_codex":0.0030506935,"about_ca_topic_score_gemma":0.002961808,"teacher_disagreement_score":0.008865261,"about_ca_system_score_codex":0.0007641707,"about_ca_system_score_gemma":0.0008308736,"threshold_uncertainty_score":0.046884596},"labels":[],"label_agreement":null},{"id":"W2026695867","doi":"10.1145/1571941.1572089","title":"Integrating phrase inseparability in phrase-based model","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Phrase; Computer science; Natural language processing; Artificial intelligence; Measure (data warehouse); Information retrieval; Data mining","score_opus":0.0304809453267647,"score_gpt":0.27919166617942504,"score_spread":0.24871072085266033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2026695867","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011001469,0.00074955233,0.9849445,0.0004342182,0.000070227194,0.000105861254,0.00016310225,0.00072580547,0.001805282],"genre_scores_gemma":[0.47532508,0.0018228613,0.5152539,0.0006947626,0.0006987801,0.0006392773,0.0014006805,0.0004162218,0.003748394],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961825,0.0018070514,0.00024347895,0.00060418603,0.0010148634,0.00014788844],"domain_scores_gemma":[0.99202394,0.0050906134,0.0004298684,0.0013557929,0.00092925766,0.00017048653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005742063,0.0012118667,0.0016903502,0.002528066,0.0008206223,0.0021995544,0.002363135,0.0015696125,0.0028320975],"category_scores_gemma":[0.01507757,0.0006644767,0.0017771858,0.0023478153,0.0010504593,0.007745174,0.0019892287,0.0034609525,0.002300107],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010904614,0.00057301816,0.0052067055,0.001069784,0.0010968021,0.0003110642,0.00077456183,0.25995308,0.042813737,0.10407743,0.009787417,0.5732459],"study_design_scores_gemma":[0.000047794434,0.00028823072,0.00089939014,0.000028358747,0.00021405071,0.00018856036,0.000040665946,0.9452936,0.004627534,0.04403607,0.0042463806,0.000089238376],"about_ca_topic_score_codex":0.0022420413,"about_ca_topic_score_gemma":0.0026182514,"teacher_disagreement_score":0.005742063,"about_ca_system_score_codex":0.00081574713,"about_ca_system_score_gemma":0.0012978882,"threshold_uncertainty_score":0.030367255},"labels":[],"label_agreement":null},{"id":"W2027974960","doi":"10.3138/infor.48.1.001","title":"An Application of Operational Research to Computational Linguistics: Word Ambiguity","year":2010,"lang":"en","type":"article","venue":"INFOR Information Systems and Operational Research","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Polysemy; Ambiguity; Computer science; Natural language processing; Word (group theory); Cluster analysis; Artificial intelligence; Measure (data warehouse); Word Association; Word lists by frequency; Principle of maximum entropy; Linguistics; Sentence; Data mining","score_opus":0.08674079423171271,"score_gpt":0.41533130772775295,"score_spread":0.32859051349604024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027974960","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003939657,0.0010797799,0.9878776,0.0019333237,0.00010641918,0.000036546757,0.00006348601,0.000080736696,0.0048824353],"genre_scores_gemma":[0.271484,0.0017706394,0.7234789,0.00075866975,0.00078111945,0.0004483534,0.00016116943,0.00019652768,0.000920546],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9859272,0.010700953,0.00061318796,0.0010331692,0.0015608731,0.00016472385],"domain_scores_gemma":[0.9394045,0.052244596,0.0019444694,0.0042869896,0.0016729963,0.00044645008],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009719498,0.0014211356,0.0016016283,0.0069067124,0.0016487223,0.003926923,0.0019415821,0.0016160196,0.0032602532],"category_scores_gemma":[0.07027993,0.00069135305,0.0017092435,0.008365824,0.012448727,0.009759312,0.0050470517,0.004138136,0.0005123031],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019037541,0.000027165439,0.0008756192,0.00016677564,0.00006371851,0.00006400646,0.0007222946,0.016521856,0.00038098404,0.93820363,0.0010108275,0.041944128],"study_design_scores_gemma":[0.000006332581,0.0000147789215,0.00020082464,0.000045015928,0.000008364055,0.00005667399,0.00015951772,0.029760659,0.00019030602,0.96547323,0.004065474,0.00001892782],"about_ca_topic_score_codex":0.0020848985,"about_ca_topic_score_gemma":0.0014482843,"teacher_disagreement_score":0.009719498,"about_ca_system_score_codex":0.0020332567,"about_ca_system_score_gemma":0.0022789433,"threshold_uncertainty_score":0.05140227},"labels":[],"label_agreement":null},{"id":"W2028279014","doi":"10.1145/2348283.2348490","title":"Lightweight contrastive summarization for news comment mining","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Theme (computing); Viewpoints; Information retrieval; Popularity; Event (particle physics); World Wide Web; Data science","score_opus":0.03445751149747339,"score_gpt":0.2630931467773359,"score_spread":0.22863563527986253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2028279014","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043689966,0.0009659239,0.9265463,0.00053613866,0.0001911953,0.0005477928,0.004781159,0.019967828,0.0027737133],"genre_scores_gemma":[0.20822404,0.00045035675,0.77011013,0.00021622503,0.0005617128,0.00073463214,0.014110312,0.0008846013,0.0047079227],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99778444,0.00059466093,0.00022172391,0.0005225453,0.00076324964,0.00011339473],"domain_scores_gemma":[0.98995435,0.00609992,0.001006232,0.0009372854,0.0017985919,0.00020356251],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024438181,0.001379787,0.0010253222,0.004595111,0.0007551605,0.0016780874,0.0018643601,0.0009998112,0.0047897906],"category_scores_gemma":[0.015149101,0.0005212379,0.000847697,0.0028818948,0.0003639806,0.0027049722,0.0015839082,0.0012168214,0.0037974366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008091861,0.00037417345,0.0072202827,0.0012477475,0.00023954982,0.0004764544,0.0013812156,0.007988174,0.09077672,0.0049195457,0.02650918,0.8580578],"study_design_scores_gemma":[0.00022851978,0.000973442,0.011167429,0.00017208402,0.00038395033,0.0011454214,0.0011991104,0.7758848,0.108153395,0.024193384,0.07629524,0.000203259],"about_ca_topic_score_codex":0.0013480624,"about_ca_topic_score_gemma":0.0032128599,"teacher_disagreement_score":0.0047897906,"about_ca_system_score_codex":0.00047428272,"about_ca_system_score_gemma":0.0007109413,"threshold_uncertainty_score":0.016023457},"labels":[],"label_agreement":null},{"id":"W2033016738","doi":"10.1145/1571941.1572087","title":"Incorporating prior knowledge into a transductive ranking algorithm for multi-document summarization","year":2009,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Automatic summarization; Ranking (information retrieval); Computer science; Artificial intelligence; Function (biology); Learning to rank; Process (computing); Information retrieval; Machine learning; Natural language processing","score_opus":0.04177516383753157,"score_gpt":0.31656684706101784,"score_spread":0.2747916832234863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033016738","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037620026,0.00016491674,0.9947798,0.00010378636,0.000020669999,0.00005822898,0.00003620108,0.0006024393,0.0004719954],"genre_scores_gemma":[0.28375244,0.00050182594,0.7087692,0.00040714093,0.00036503794,0.0006021217,0.00084570504,0.00046768677,0.004288902],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997755,0.0010500755,0.00014080867,0.00042093184,0.0005293717,0.00010383897],"domain_scores_gemma":[0.99462974,0.0037491403,0.00028450546,0.0004334088,0.0008143305,0.00008893281],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042122165,0.0014944316,0.0019431897,0.0021586912,0.00066742767,0.0015926136,0.002160187,0.0017102106,0.0025814406],"category_scores_gemma":[0.009912307,0.0007335994,0.0013217259,0.0016113764,0.00084638014,0.0036234355,0.0012233957,0.0025900921,0.0014950566],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017536547,0.0003638697,0.0007598665,0.00037424645,0.00021380548,0.000105759966,0.0003740331,0.4229951,0.011854463,0.021877153,0.004807833,0.5360985],"study_design_scores_gemma":[0.000010388546,0.00009600695,0.00012800419,0.000012283511,0.000025711946,0.000023015005,0.000017804114,0.982659,0.001814897,0.014335408,0.00086112285,0.000016472475],"about_ca_topic_score_codex":0.0012962458,"about_ca_topic_score_gemma":0.002613018,"teacher_disagreement_score":0.0042122165,"about_ca_system_score_codex":0.0012548999,"about_ca_system_score_gemma":0.0007860321,"threshold_uncertainty_score":0.02227658},"labels":[],"label_agreement":null},{"id":"W2033839684","doi":"10.3115/1613715.1613813","title":"Summarizing spoken and written conversations","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":91,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Conversation; Computer science; Domain (mathematical analysis); Natural language processing; Open domain; Artificial intelligence; Information retrieval; World Wide Web; Linguistics; Question answering","score_opus":0.03376233065434065,"score_gpt":0.22319047416338736,"score_spread":0.18942814350904671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033839684","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08515335,0.0062447456,0.88160414,0.0011843067,0.00095198886,0.0007515071,0.004709775,0.007316843,0.012083435],"genre_scores_gemma":[0.43077725,0.0029909895,0.5367715,0.000489292,0.0014463963,0.0005854177,0.016021997,0.00096733775,0.009949766],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963595,0.0014540759,0.00026923468,0.0009959429,0.0007407672,0.00018043543],"domain_scores_gemma":[0.9924609,0.0039763493,0.0006498768,0.0011046908,0.0016146103,0.00019363011],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002363106,0.0016748051,0.0010488244,0.0028015985,0.0011009234,0.0027401627,0.0011155636,0.0009292539,0.0047004065],"category_scores_gemma":[0.015065515,0.00043639366,0.00095568364,0.002137362,0.00039396522,0.0036037518,0.0019126557,0.0012969582,0.002309917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008452965,0.00017109582,0.0044234986,0.0018227496,0.0004152836,0.00031401915,0.0042370134,0.013368284,0.06689553,0.007575722,0.016975934,0.8829555],"study_design_scores_gemma":[0.00017615758,0.0011688374,0.028310569,0.00097469834,0.0017790243,0.0010810464,0.011272639,0.49585834,0.1448136,0.084984794,0.2290794,0.0005009431],"about_ca_topic_score_codex":0.0018693979,"about_ca_topic_score_gemma":0.0027837693,"teacher_disagreement_score":0.0047004065,"about_ca_system_score_codex":0.00052625354,"about_ca_system_score_gemma":0.0009143253,"threshold_uncertainty_score":0.015724361},"labels":[],"label_agreement":null},{"id":"W2038336484","doi":"10.1145/2492517.2500308","title":"A system for the automated author attribution of text and instant messages","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Instant; Instant messaging; Authorship attribution; Attribution; Computer science; Supervisor; Artificial intelligence; Naive Bayes classifier; Multimedia; World Wide Web; Natural language processing; Psychology; Social psychology","score_opus":0.030444102845373716,"score_gpt":0.26435391313543793,"score_spread":0.23390981029006422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2038336484","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014083343,0.00045446874,0.65048766,0.00050618267,0.0005661913,0.0007249391,0.007400224,0.31978253,0.005994607],"genre_scores_gemma":[0.17603348,0.000503655,0.7542746,0.0006545493,0.0011001208,0.0013014338,0.028849652,0.010242238,0.027040241],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99392045,0.0013453468,0.00068937277,0.001530367,0.0023118828,0.00020266189],"domain_scores_gemma":[0.9804007,0.0065918406,0.0021545622,0.0049624187,0.0044822353,0.0014083765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00522657,0.001432963,0.0013294765,0.0074872365,0.0015437051,0.0023142155,0.0018497109,0.0019051905,0.011642543],"category_scores_gemma":[0.021037763,0.0009800787,0.0008154841,0.0029317902,0.0005367156,0.005883316,0.0038033882,0.0016504516,0.018128948],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010430651,0.00040809577,0.01265859,0.0006655764,0.00018068354,0.00054911524,0.0015413181,0.0035243358,0.030431021,0.007958268,0.20508747,0.73595256],"study_design_scores_gemma":[0.00034286766,0.000709712,0.02144008,0.00034332997,0.0002755546,0.002273601,0.000806971,0.45858723,0.10699101,0.040566172,0.36694685,0.0007165865],"about_ca_topic_score_codex":0.0024449527,"about_ca_topic_score_gemma":0.0027448027,"teacher_disagreement_score":0.011642543,"about_ca_system_score_codex":0.00096455496,"about_ca_system_score_gemma":0.0019777808,"threshold_uncertainty_score":0.03894818},"labels":[],"label_agreement":null},{"id":"W2040916592","doi":"10.1145/2661829.2661887","title":"Robust Entity Linking via Random Walks","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates - Technology Futures","keywords":"Computer science; Benchmark (surveying); Random walk; Information retrieval; Representation (politics); Context (archaeology); Knowledge base; Entity linking; Artificial intelligence; Popularity; Semantic similarity; Feature (linguistics); Base (topology); Similarity (geometry); Natural language processing; Task (project management)","score_opus":0.024636894154808155,"score_gpt":0.2140333906751774,"score_spread":0.18939649652036925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2040916592","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024555286,0.0007857798,0.9654885,0.00027344245,0.00006385981,0.00017703735,0.00068087765,0.005985387,0.0019898075],"genre_scores_gemma":[0.36318982,0.00082958746,0.61420584,0.00045528013,0.0002449493,0.0003636819,0.0079561975,0.0013460693,0.011408616],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99707615,0.00083753286,0.00013814081,0.0011407132,0.00063562713,0.00017182558],"domain_scores_gemma":[0.99328893,0.004096371,0.00068668684,0.0012482376,0.00051665306,0.00016314801],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022130704,0.0018972838,0.0022489831,0.006057114,0.0013953116,0.0024515255,0.0034735366,0.0035426298,0.0034729177],"category_scores_gemma":[0.012773768,0.0012682491,0.0017012666,0.0069948803,0.0010696617,0.005647786,0.003199953,0.001910264,0.003695143],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032542984,0.00028231856,0.0037624226,0.0003196853,0.0003050701,0.00039984105,0.0002136903,0.59653616,0.005948596,0.02435079,0.014922148,0.3526339],"study_design_scores_gemma":[0.000019045247,0.000027966442,0.00025451632,0.00001435509,0.000028370623,0.00012117278,0.000025011095,0.9734355,0.0020723017,0.02154708,0.002437837,0.000016710852],"about_ca_topic_score_codex":0.00489413,"about_ca_topic_score_gemma":0.0076993513,"teacher_disagreement_score":0.006057114,"about_ca_system_score_codex":0.00083905767,"about_ca_system_score_gemma":0.0012442268,"threshold_uncertainty_score":0.011704028},"labels":[],"label_agreement":null},{"id":"W2042408352","doi":"10.3115/1614108.1614144","title":"Stating with certainty or stating with doubt","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Certainty; Statement (logic); Computer science; Focus (optics); Annotation; Perspective (graphical); Psychology; Epistemology; Artificial intelligence; Philosophy","score_opus":0.025404467894124275,"score_gpt":0.26501203491369296,"score_spread":0.2396075670195687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2042408352","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6049598,0.0037767347,0.25954068,0.00890594,0.0010671922,0.00041595657,0.0010688026,0.0008236318,0.11944122],"genre_scores_gemma":[0.9858272,0.00036012643,0.011505532,0.0004068505,0.00017801873,0.0000761557,0.00017296738,0.000055412384,0.0014176324],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98891467,0.0055372682,0.0013018519,0.0013531368,0.002439646,0.0004534358],"domain_scores_gemma":[0.9422268,0.037341084,0.0112392595,0.004607887,0.0037834216,0.0008014038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073621552,0.000646219,0.00038839644,0.0012561345,0.0014344335,0.003193524,0.0007251795,0.0012408604,0.003230322],"category_scores_gemma":[0.061049547,0.00029125396,0.00043023855,0.0010210259,0.005362177,0.006025703,0.0034206407,0.0017856234,0.00052567077],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015284927,0.00008422335,0.036267344,0.002628182,0.0001977133,0.001535396,0.18487002,0.0021354076,0.028434642,0.55737835,0.007983847,0.17695633],"study_design_scores_gemma":[0.000080563244,0.0005675795,0.052827872,0.0019814968,0.0003829839,0.0051414818,0.08967781,0.013596208,0.02563523,0.6186179,0.19100451,0.00048634765],"about_ca_topic_score_codex":0.00064946763,"about_ca_topic_score_gemma":0.00057018985,"teacher_disagreement_score":0.0073621552,"about_ca_system_score_codex":0.00086769403,"about_ca_system_score_gemma":0.00081163854,"threshold_uncertainty_score":0.038935244},"labels":[],"label_agreement":null},{"id":"W2045517286","doi":"10.1109/icdim.2010.5664669","title":"Latent semantic indexing and large dataset: Study of term-weighting schemes","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"","keywords":"Computer science; Weighting; Information retrieval; Term (time); Search engine indexing; Precision and recall; Document retrieval; Latent semantic analysis; Graph; Term Discrimination; Data mining; Search engine; Concept search; Web search query","score_opus":0.02180303817818376,"score_gpt":0.2713560739859011,"score_spread":0.24955303580771734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2045517286","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6344007,0.010134133,0.33206126,0.0031761609,0.0004802732,0.0011578939,0.012140956,0.0022445759,0.004203959],"genre_scores_gemma":[0.82603765,0.0014159532,0.14988773,0.0003011807,0.0003349113,0.00066159456,0.020043178,0.00020581445,0.0011119074],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9820219,0.009316021,0.0016484044,0.0026418436,0.0038325563,0.00053929235],"domain_scores_gemma":[0.8768646,0.09906382,0.0054333713,0.012256945,0.005474999,0.0009062686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034030374,0.0010918883,0.0014428206,0.006456625,0.0016660143,0.0030325123,0.0021022884,0.0021106023,0.0007040037],"category_scores_gemma":[0.10443954,0.00030706255,0.0016116635,0.009365052,0.0012905599,0.006540039,0.0019273217,0.0022858374,0.00028191364],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029264097,0.0028258087,0.21107924,0.0023073566,0.0028723816,0.0005215782,0.0013494086,0.20175196,0.014108812,0.04188096,0.025388265,0.49298778],"study_design_scores_gemma":[0.00023894628,0.0010177481,0.055341907,0.00017855351,0.00043166586,0.0008387151,0.00072878005,0.8968669,0.008258031,0.027129518,0.008768452,0.00020065109],"about_ca_topic_score_codex":0.007918057,"about_ca_topic_score_gemma":0.0060420185,"teacher_disagreement_score":0.034030374,"about_ca_system_score_codex":0.0024421585,"about_ca_system_score_gemma":0.0013432162,"threshold_uncertainty_score":0.179972},"labels":[],"label_agreement":null},{"id":"W2045929671","doi":"10.1145/1376815.1376819","title":"Semantic text similarity using corpus-based word similarity and string similarity","year":2008,"lang":"en","type":"article","venue":"ACM Transactions on Knowledge Discovery from Data","topic":"Topic Modeling","field":"Computer Science","cited_by":485,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Similarity (geometry); Semantic similarity; String metric; Artificial intelligence; Word (group theory); Natural language processing; Longest common subsequence problem; Focus (optics); String searching algorithm; String (physics); Representation (politics); Information retrieval; Matching (statistics); Pattern matching; Mathematics; Algorithm","score_opus":0.12043376955524356,"score_gpt":0.3074205502158795,"score_spread":0.18698678066063595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2045929671","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051913235,0.0021264022,0.930313,0.00026806977,0.0004651728,0.0006221626,0.0023097303,0.0025874637,0.009394729],"genre_scores_gemma":[0.2841628,0.0011061836,0.7038279,0.00012819398,0.0005406998,0.0010794784,0.00547775,0.00041786322,0.0032591443],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99439275,0.0011248859,0.00070041994,0.0010071052,0.0026378462,0.0001368597],"domain_scores_gemma":[0.99211943,0.0030306864,0.0009637382,0.0011390698,0.0024947596,0.00025224863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021779973,0.00092376,0.0014646518,0.017746313,0.0009860167,0.002976211,0.0016307629,0.0012851419,0.003514882],"category_scores_gemma":[0.018149171,0.00031206195,0.0011634456,0.016821468,0.0012409195,0.005136708,0.0020246366,0.0010016097,0.0019029668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005017755,0.00040544325,0.012898957,0.001379963,0.00056546653,0.00041738807,0.0011314675,0.02129064,0.028761242,0.053923283,0.011165284,0.86755913],"study_design_scores_gemma":[0.00019402873,0.00092868815,0.027077055,0.00043689326,0.00055513874,0.0028144156,0.0016984984,0.67202175,0.048917342,0.18239895,0.062490948,0.0004662709],"about_ca_topic_score_codex":0.0016587181,"about_ca_topic_score_gemma":0.0016378546,"teacher_disagreement_score":0.017746313,"about_ca_system_score_codex":0.000980401,"about_ca_system_score_gemma":0.0013563224,"threshold_uncertainty_score":0.011758447},"labels":[],"label_agreement":null},{"id":"W2046353967","doi":"10.1145/1740592.1740594","title":"Exploiting query logs for cross-lingual query suggestions","year":2010,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Microsoft Research Asia; Chinese University of Hong Kong; Innovation and Technology Commission","keywords":"Computer science; Cross-language information retrieval; Query expansion; Query language; RDF query language; Query optimization; Information retrieval; Web query classification; Sargable; Natural language processing; Relevance (law); Web search query; Discriminative model; Query by Example; Artificial intelligence; Search engine","score_opus":0.027786355921336915,"score_gpt":0.2897494147459037,"score_spread":0.2619630588245668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046353967","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15674108,0.005383803,0.7637402,0.0025693239,0.0004349156,0.0016486059,0.0069594323,0.05466308,0.007859611],"genre_scores_gemma":[0.69882447,0.0013162754,0.27968967,0.0005582467,0.0002844184,0.00073734793,0.014046262,0.001414001,0.0031293563],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9915353,0.0038295118,0.0007176275,0.0010733542,0.0025169877,0.0003272876],"domain_scores_gemma":[0.95759994,0.027358532,0.0022227785,0.005802633,0.0063659055,0.000650103],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007814266,0.002422083,0.002437064,0.0055109696,0.0012460272,0.003119016,0.0022491768,0.0013263121,0.0034122667],"category_scores_gemma":[0.042332273,0.0012698016,0.00091907807,0.004314958,0.0006778713,0.0091549,0.0025852274,0.002520194,0.003177269],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039820955,0.0020892294,0.030157465,0.002504461,0.0004433674,0.0011182622,0.0029673895,0.036595244,0.07507897,0.006487144,0.04029147,0.7982849],"study_design_scores_gemma":[0.00024969876,0.0007000424,0.008089249,0.00012984654,0.00027433565,0.00073127775,0.0011029857,0.9113534,0.03740725,0.012850291,0.02686367,0.0002479376],"about_ca_topic_score_codex":0.008908823,"about_ca_topic_score_gemma":0.012763482,"teacher_disagreement_score":0.008908823,"about_ca_system_score_codex":0.0010832853,"about_ca_system_score_gemma":0.0031007638,"threshold_uncertainty_score":0.041326344},"labels":[],"label_agreement":null},{"id":"W2046573508","doi":"10.3758/brm.40.3.705","title":"WINDSOR: Windsor improved norms of distance and similarity of representations of semantics","year":2008,"lang":"en","type":"article","venue":"Behavior Research Methods","topic":"Topic Modeling","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Windsor; Computer science; Similarity (geometry); Word (group theory); Semantic similarity; Semantics (computer science); Meaning (existential); Natural language processing; Lexical semantics; Cognition; Artificial intelligence; Lexical item; Linguistics; Psychology","score_opus":0.27716463127891805,"score_gpt":0.5285698993187112,"score_spread":0.2514052680397932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046573508","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014935969,0.00029253305,0.97468036,0.00026404552,0.00022709742,0.00016015202,0.0015513867,0.0046478244,0.0032406887],"genre_scores_gemma":[0.13674758,0.00031725428,0.8463254,0.0001439774,0.00020729404,0.00059425156,0.004571745,0.002953595,0.008138926],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935127,0.0022519238,0.0005458059,0.0012885773,0.0021694938,0.00023160805],"domain_scores_gemma":[0.988882,0.0041603097,0.0006012839,0.002549015,0.0033003415,0.00050696725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063635074,0.0013519222,0.0014844524,0.004225905,0.0011421286,0.003657533,0.0022319697,0.0014756564,0.008322127],"category_scores_gemma":[0.032628182,0.0008678781,0.0016044735,0.0034037416,0.0012843221,0.008093067,0.004477959,0.0026557331,0.0027402982],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015234284,0.00039660966,0.005001825,0.0006043115,0.00036845877,0.00010554678,0.0016423987,0.039946005,0.008601024,0.26257667,0.039202902,0.6400309],"study_design_scores_gemma":[0.0003703688,0.00047917906,0.0029857892,0.00017054401,0.00019083598,0.0002330642,0.0005489913,0.56727374,0.016136117,0.33271298,0.0787107,0.00018773835],"about_ca_topic_score_codex":0.008511922,"about_ca_topic_score_gemma":0.013523895,"teacher_disagreement_score":0.008511922,"about_ca_system_score_codex":0.0015698277,"about_ca_system_score_gemma":0.0020104914,"threshold_uncertainty_score":0.033653855},"labels":[],"label_agreement":null},{"id":"W2047469196","doi":"10.1037/h0087384","title":"Strategies of text retrieval: A criterion shift account.","year":2002,"lang":"en","type":"article","venue":"Canadian Journal of Experimental Psychology/Revue canadienne de psychologie expérimentale","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Psychology; Paraphrase; Set (abstract data type); Test (biology); Inference; Statistics; Natural language processing; Social psychology; Cognitive psychology; Artificial intelligence; Information retrieval; Computer science; Mathematics","score_opus":0.05658541191767303,"score_gpt":0.307067143894818,"score_spread":0.25048173197714496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2047469196","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89646965,0.00070865836,0.09031018,0.00069674005,0.000026009551,0.00025998478,0.00008486995,0.00031779724,0.011126163],"genre_scores_gemma":[0.98264277,0.00011410589,0.016114246,0.00009833165,0.000018812869,0.000101222744,0.00007442515,0.000038541213,0.0007974746],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9951839,0.0018099374,0.00028333225,0.0009428547,0.001542483,0.00023751492],"domain_scores_gemma":[0.9603423,0.026001075,0.0046894304,0.004216289,0.0036950237,0.0010559071],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073049245,0.0006477152,0.00051328406,0.0016302009,0.000514249,0.0025171954,0.0013608737,0.001634679,0.0017275901],"category_scores_gemma":[0.077197894,0.00041020024,0.00056844583,0.0008286613,0.0015259824,0.003971555,0.0018286313,0.00086572155,0.0006338511],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025606395,0.0011309562,0.2174023,0.0012453939,0.00078807794,0.001184332,0.030995635,0.01534047,0.18536904,0.099598445,0.0023355123,0.44204918],"study_design_scores_gemma":[0.0009241585,0.0040226816,0.3047626,0.00023055759,0.00069063104,0.0039163777,0.00886684,0.29422298,0.072823524,0.29804766,0.01104841,0.00044359502],"about_ca_topic_score_codex":0.0010130903,"about_ca_topic_score_gemma":0.00083218963,"teacher_disagreement_score":0.0073049245,"about_ca_system_score_codex":0.000742872,"about_ca_system_score_gemma":0.00072779274,"threshold_uncertainty_score":0.03863257},"labels":[],"label_agreement":null},{"id":"W2048614195","doi":"10.1145/2641483.2641534","title":"Identifying Questions &amp; Requests in Conversation","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Interrogative; Conversation; Computer science; Sentence; Natural language processing; Converse; Punctuation; Artificial intelligence; Spoken language; Class (philosophy); Phrase; Interrogative word; Linguistics","score_opus":0.08462228065266532,"score_gpt":0.3026187454326794,"score_spread":0.2179964647800141,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2048614195","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8664066,0.0015095755,0.09915185,0.001737168,0.0002329222,0.00062910264,0.008779589,0.0064679235,0.015085186],"genre_scores_gemma":[0.93228376,0.0003497354,0.048444945,0.00032132727,0.00019692993,0.00031014683,0.012058141,0.00023498258,0.005800094],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9945509,0.0028180387,0.00034233963,0.0010023769,0.00079678366,0.00048960204],"domain_scores_gemma":[0.9917378,0.0048722927,0.0008280511,0.0005811354,0.0014882201,0.00049249677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023392786,0.0009149593,0.00074585446,0.0024868844,0.0013953662,0.0018801965,0.0009062868,0.0017262313,0.0024071534],"category_scores_gemma":[0.011254139,0.0003882524,0.0006857805,0.0013010071,0.00043846734,0.0032065471,0.0015425244,0.0009404801,0.0034751347],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034430407,0.0011981352,0.20946455,0.0022052794,0.00027336102,0.0039286744,0.03400285,0.010801033,0.12225901,0.010044658,0.08852403,0.51385534],"study_design_scores_gemma":[0.00015834566,0.0007223013,0.21949385,0.0002759436,0.00030739437,0.0046303794,0.036440548,0.4800503,0.090108074,0.0186457,0.14879888,0.00036836628],"about_ca_topic_score_codex":0.007827057,"about_ca_topic_score_gemma":0.0091733625,"teacher_disagreement_score":0.007827057,"about_ca_system_score_codex":0.0010786909,"about_ca_system_score_gemma":0.0009039914,"threshold_uncertainty_score":0.015562952},"labels":[],"label_agreement":null},{"id":"W2049557239","doi":"10.1016/j.eswa.2015.04.054","title":"Evolutionary fine-tuning of automated semantic annotation systems","year":2015,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Annotation; Task (project management); Process (computing); Artificial intelligence; Natural language processing; Domain (mathematical analysis); Information retrieval; Programming language","score_opus":0.030865148638467458,"score_gpt":0.2729550098003744,"score_spread":0.24208986116190695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049557239","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19088884,0.00078521395,0.79340655,0.000800975,0.0002530551,0.00034424124,0.00036431872,0.0060957153,0.0070610684],"genre_scores_gemma":[0.6925906,0.00019956013,0.30103222,0.00024843265,0.00009721352,0.00024339368,0.0011390548,0.0011870502,0.0032625303],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995337,0.0020936285,0.00025874036,0.0011875356,0.00070471765,0.000418383],"domain_scores_gemma":[0.983567,0.010565107,0.00047858525,0.0024424475,0.002570083,0.00037671314],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005662939,0.0010375615,0.0016033143,0.0025801014,0.0013236874,0.0024435252,0.0028377818,0.0022018505,0.0041840873],"category_scores_gemma":[0.026621241,0.0009193548,0.0011892788,0.0016525058,0.0010255651,0.0030570782,0.0030166267,0.0019752795,0.0014342366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064550724,0.0005697532,0.009303406,0.0003657942,0.0002658873,0.00025778793,0.0009772631,0.31638095,0.028713988,0.018710546,0.008748446,0.61506075],"study_design_scores_gemma":[0.00003455079,0.00004794835,0.00064553536,0.000016380798,0.000057750163,0.0000521402,0.00016074454,0.98166496,0.003961576,0.011568923,0.0017767196,0.000012783135],"about_ca_topic_score_codex":0.0046572424,"about_ca_topic_score_gemma":0.007932248,"teacher_disagreement_score":0.005662939,"about_ca_system_score_codex":0.0018244784,"about_ca_system_score_gemma":0.0023032024,"threshold_uncertainty_score":0.02994883},"labels":[],"label_agreement":null},{"id":"W2050611670","doi":"10.1371/journal.pcbi.1000391","title":"How to Get the Most out of Your Curation Effort","year":2009,"lang":"en","type":"article","venue":"PLoS Computational Biology","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"U.S. National Library of Medicine; National Institute of General Medical Sciences; National Institutes of Health; National Science Foundation","keywords":"Annotation; Computer science; Data curation; Sentence; Probabilistic logic; Set (abstract data type); Natural language processing; Task (project management); Information retrieval; Artificial intelligence; Machine learning; Data science; Data mining","score_opus":0.046541702477658514,"score_gpt":0.2859339761740192,"score_spread":0.2393922736963607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050611670","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039034538,0.0042534703,0.8849812,0.037174463,0.0011694351,0.0005760849,0.002726512,0.012330589,0.017753651],"genre_scores_gemma":[0.20282248,0.0021445113,0.7778029,0.0035006716,0.000564202,0.00052999746,0.0022085912,0.0038616115,0.0065650214],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9749274,0.011934867,0.0015879167,0.00570954,0.0049872305,0.00085297495],"domain_scores_gemma":[0.9275143,0.034561016,0.0057191937,0.01575757,0.01418416,0.0022638724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028748635,0.0019639782,0.0022705551,0.004093243,0.0030573618,0.008295806,0.0029083719,0.0036696503,0.0060200826],"category_scores_gemma":[0.12673736,0.001396306,0.001725226,0.0041159783,0.0021613757,0.0131307775,0.0035841882,0.0026483336,0.009386887],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006369033,0.00025273752,0.018119926,0.0021758578,0.0005438358,0.00041021974,0.0068390933,0.0125070615,0.019397564,0.023436137,0.16611405,0.7495665],"study_design_scores_gemma":[0.0003319077,0.0005378872,0.022026569,0.001945483,0.00076621806,0.0021611194,0.011009302,0.15326425,0.045290712,0.29978094,0.46162575,0.0012598439],"about_ca_topic_score_codex":0.00499262,"about_ca_topic_score_gemma":0.008634796,"teacher_disagreement_score":0.028748635,"about_ca_system_score_codex":0.001617573,"about_ca_system_score_gemma":0.0044125863,"threshold_uncertainty_score":0.15203911},"labels":[],"label_agreement":null},{"id":"W2050715302","doi":"10.1145/1982185.1982243","title":"Multi-document summarization of scientific corpora","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Automatic summarization; Computer science; Vocabulary; Pyramid (geometry); Natural language processing; Multi-document summarization; Information retrieval; Sentence; Artificial intelligence; Focus (optics); Linguistics; Mathematics","score_opus":0.0767327168714705,"score_gpt":0.2505326633670444,"score_spread":0.1737999464955739,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050715302","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09144285,0.0067641293,0.8756065,0.0007583604,0.00047782485,0.0012784244,0.004503871,0.013793942,0.005374038],"genre_scores_gemma":[0.18241775,0.0015509633,0.79619646,0.000121986894,0.00030782685,0.0006072298,0.013087203,0.0007227584,0.0049877847],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99746776,0.00093022664,0.00031647045,0.0005415961,0.0006441253,0.00009988954],"domain_scores_gemma":[0.9926662,0.0029921453,0.0006914587,0.0009155274,0.0025789316,0.00015584928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031551744,0.0015334236,0.001323541,0.0056679808,0.000798068,0.0018042902,0.0010438736,0.0006577496,0.0022083088],"category_scores_gemma":[0.012017226,0.00033456966,0.000981609,0.0038738544,0.0002716069,0.0031547903,0.00092709955,0.0007577506,0.0014673965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036081555,0.0001790659,0.0018131763,0.0019274398,0.00028160276,0.00021905346,0.0008916561,0.020375295,0.04326591,0.0028672575,0.010765909,0.91705275],"study_design_scores_gemma":[0.00035207564,0.002058532,0.01569016,0.00041038732,0.0015084319,0.0009777794,0.0035669233,0.56537646,0.25833213,0.022451274,0.12897795,0.00029785102],"about_ca_topic_score_codex":0.001595515,"about_ca_topic_score_gemma":0.003768112,"teacher_disagreement_score":0.0056679808,"about_ca_system_score_codex":0.0007044356,"about_ca_system_score_gemma":0.0008793934,"threshold_uncertainty_score":0.01668632},"labels":[],"label_agreement":null},{"id":"W2056469463","doi":"10.3115/1072228.1072376","title":"Investigating the relationship between word segmentation performance and retrieval performance in Chinese IR","year":2002,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Segmentation; Text segmentation; Artificial intelligence; Natural language processing; Word (group theory); Information retrieval; Pattern recognition (psychology); Linguistics","score_opus":0.06429135091460345,"score_gpt":0.26874737942535387,"score_spread":0.20445602851075043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056469463","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99561834,0.0003639386,0.002231304,0.00008574085,0.0000053377735,0.00001712057,0.0001173765,0.00008011415,0.0014806986],"genre_scores_gemma":[0.99847394,0.00008381236,0.00083753635,0.000016620485,0.000014208608,0.0000097223065,0.00030903396,0.000025673757,0.00022943213],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969627,0.0012082371,0.0003480456,0.0005852898,0.0005273358,0.00036834902],"domain_scores_gemma":[0.91333723,0.071226306,0.006703774,0.0037075602,0.0041519105,0.00087319996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076307585,0.00053961435,0.0005521141,0.001848836,0.00064440654,0.0016916051,0.00045903606,0.0008261612,0.001107147],"category_scores_gemma":[0.04714511,0.0002872709,0.00042275887,0.0035076928,0.0011230899,0.0032963036,0.00071628054,0.00083553366,0.00072004955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020491066,0.0003945624,0.78713924,0.000596009,0.0005268434,0.0003407932,0.0050554005,0.037291862,0.037081763,0.0014964258,0.0015192723,0.12650873],"study_design_scores_gemma":[0.000026065622,0.0010907255,0.8935361,0.000019439964,0.00019244391,0.00023341186,0.0012756321,0.084618345,0.016984776,0.001259482,0.00066599285,0.000097656964],"about_ca_topic_score_codex":0.009180247,"about_ca_topic_score_gemma":0.0067235655,"teacher_disagreement_score":0.009180247,"about_ca_system_score_codex":0.000865082,"about_ca_system_score_gemma":0.00063333364,"threshold_uncertainty_score":0.0403558},"labels":[],"label_agreement":null},{"id":"W2060686","doi":"10.63317/47puorupsub8","title":"Automatic Summarization Using Terminological and Semantic Resources","year":2010,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Automatic summarization; Computer science; Natural language processing; Information retrieval; Artificial intelligence","score_opus":0.04781014630795609,"score_gpt":0.2754713633292548,"score_spread":0.2276612170212987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060686","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011831849,0.0012782954,0.9782245,0.00023802779,0.00008194865,0.00015413527,0.0010308011,0.0057276143,0.0014327943],"genre_scores_gemma":[0.08477559,0.0010893875,0.9045154,0.00008085212,0.00023489325,0.00026728303,0.0066569713,0.0005722985,0.0018073035],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980039,0.0006204105,0.00021797199,0.0004145346,0.0006685405,0.00007474959],"domain_scores_gemma":[0.9959229,0.0017264251,0.00054336525,0.00046507485,0.0012708883,0.00007128489],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022318827,0.0013133191,0.0012341923,0.0074197883,0.00077862374,0.0019900003,0.0011939768,0.0008423128,0.0021558048],"category_scores_gemma":[0.007826054,0.00044886637,0.0010492302,0.0037072368,0.00041881224,0.0037049444,0.0007958793,0.0009097439,0.0020033466],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031752104,0.0001381296,0.0018194644,0.001210296,0.00027265857,0.0002781734,0.00065088266,0.017937293,0.05959615,0.016250622,0.012395665,0.8891331],"study_design_scores_gemma":[0.0001511952,0.00047789473,0.006933155,0.0004389518,0.0008857041,0.0008685611,0.0010087551,0.64240444,0.13773611,0.09782881,0.11101454,0.00025194764],"about_ca_topic_score_codex":0.0009222443,"about_ca_topic_score_gemma":0.0014774778,"teacher_disagreement_score":0.0074197883,"about_ca_system_score_codex":0.00069030595,"about_ca_system_score_gemma":0.0007885033,"threshold_uncertainty_score":0.011803448},"labels":[],"label_agreement":null},{"id":"W2060733617","doi":"10.3115/1708155.1708159","title":"Optimization-based content selection for opinion summarization","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Toronto","funders":"","keywords":"Automatic summarization; Selection (genetic algorithm); Computer science; Content (measure theory); Cluster analysis; Realization (probability); Heuristic; Multi-document summarization; Data mining; Information retrieval; Artificial intelligence; Mathematics; Statistics","score_opus":0.05149195832567969,"score_gpt":0.266172036641591,"score_spread":0.21468007831591132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2060733617","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042440253,0.00015572568,0.9934807,0.00013762338,0.000036240555,0.0001268187,0.000101345875,0.0006403433,0.001077184],"genre_scores_gemma":[0.15795581,0.00021509873,0.83822316,0.00014932513,0.00021506652,0.00053620286,0.0007780489,0.00023198285,0.0016953308],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971752,0.0013131297,0.00016138612,0.00047561605,0.0007558554,0.000118713266],"domain_scores_gemma":[0.995669,0.0023961524,0.00041504644,0.0003584259,0.0010564241,0.00010498962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002711034,0.001131468,0.0013187691,0.0027414176,0.0008935279,0.0015576035,0.0015395698,0.0010848633,0.003408084],"category_scores_gemma":[0.011706469,0.00035727667,0.0008747489,0.0020952248,0.0006610718,0.0019702427,0.001080097,0.0010162562,0.001537106],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035496827,0.00022521117,0.001605816,0.0006415296,0.00021687822,0.0001397027,0.0006664379,0.1266034,0.028786905,0.034804303,0.017438661,0.78851616],"study_design_scores_gemma":[0.00006789878,0.00010737539,0.00068255956,0.000035103556,0.00006808338,0.00007015875,0.00015616517,0.9485829,0.012827709,0.029389866,0.007971187,0.00004085059],"about_ca_topic_score_codex":0.0015384284,"about_ca_topic_score_gemma":0.0019738842,"teacher_disagreement_score":0.003408084,"about_ca_system_score_codex":0.0012616725,"about_ca_system_score_gemma":0.001063671,"threshold_uncertainty_score":0.01433748},"labels":[],"label_agreement":null},{"id":"W2061282927","doi":"10.1145/2494266.2494280","title":"A graph-based topic extraction method enabling simple interactive customization","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs; Boeing","keywords":"Computer science; Interpretability; Latent Dirichlet allocation; Cluster analysis; Personalization; Artificial intelligence; Graph; Data mining; Machine learning; Representation (politics); Matrix decomposition; Set (abstract data type); Topic model; Information retrieval; Theoretical computer science","score_opus":0.02366903509246583,"score_gpt":0.3098053513791356,"score_spread":0.2861363162866698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061282927","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033728457,0.00011234084,0.9919853,0.00008329241,0.00003207248,0.00010923122,0.00027861746,0.0034477713,0.00057851354],"genre_scores_gemma":[0.040276293,0.00020086559,0.95502794,0.00005797532,0.00005308215,0.0002986032,0.0014259438,0.0005843279,0.0020749767],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99882275,0.00030922832,0.00009267759,0.00038085933,0.00032956223,0.00006497262],"domain_scores_gemma":[0.9981244,0.00093718205,0.00013741798,0.00029831752,0.00043940134,0.000063195235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011266163,0.0010437106,0.00082157494,0.0043031904,0.00094953005,0.0014355349,0.00095688493,0.00092456356,0.0042541586],"category_scores_gemma":[0.0045422744,0.0006383069,0.0013657562,0.0037746208,0.0005821269,0.002020094,0.0011990818,0.0011690853,0.002949996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003234361,0.00013338446,0.0014967248,0.00049563363,0.00015656883,0.00029895155,0.0007728823,0.021205306,0.07412442,0.018858379,0.018706733,0.86342764],"study_design_scores_gemma":[0.00015182162,0.00014978305,0.0041439435,0.000100696525,0.00019673858,0.0015584157,0.00046298146,0.7765829,0.06573034,0.05388024,0.09683139,0.00021086569],"about_ca_topic_score_codex":0.002762221,"about_ca_topic_score_gemma":0.004214411,"teacher_disagreement_score":0.0043031904,"about_ca_system_score_codex":0.0005154843,"about_ca_system_score_gemma":0.0010520833,"threshold_uncertainty_score":0.014231563},"labels":[],"label_agreement":null},{"id":"W2065354331","doi":"10.3115/1708322.1708330","title":"Extractive vs. NLG-based abstractive summarization of evaluative text","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Natural language processing; Margin (machine learning); Context (archaeology); Artificial intelligence; Abstraction; Natural language generation; Information retrieval; Natural language; Machine learning","score_opus":0.041963495997206106,"score_gpt":0.28514386411850234,"score_spread":0.24318036812129623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2065354331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42277777,0.0029840039,0.54375756,0.0013175174,0.0003178481,0.0011656182,0.004611932,0.011482247,0.011585499],"genre_scores_gemma":[0.6649174,0.0010448484,0.32095683,0.00022844298,0.0002697444,0.00051854184,0.007932482,0.0007485822,0.0033831412],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99711037,0.0014574515,0.00035439586,0.00039247226,0.0005887033,0.000096537755],"domain_scores_gemma":[0.98454434,0.009636166,0.0016973938,0.0021256828,0.0017981939,0.00019831593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003978954,0.0009857228,0.0009047575,0.0022546796,0.00038057924,0.0018583703,0.000912495,0.00063055404,0.001997974],"category_scores_gemma":[0.020945614,0.00028521157,0.000658287,0.0021409488,0.0004648695,0.0027435708,0.0010803291,0.00088962977,0.0013691697],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026883576,0.0003998904,0.008153856,0.0019901553,0.00047294883,0.00024692953,0.0030809522,0.044283323,0.1405177,0.005233876,0.008731929,0.7842],"study_design_scores_gemma":[0.0004948563,0.0028040423,0.040596638,0.00024090795,0.0011049821,0.0007398271,0.0025415183,0.65921474,0.24309967,0.017220201,0.03163673,0.00030584825],"about_ca_topic_score_codex":0.00064238725,"about_ca_topic_score_gemma":0.0013805026,"teacher_disagreement_score":0.003978954,"about_ca_system_score_codex":0.00043749827,"about_ca_system_score_gemma":0.00042385707,"threshold_uncertainty_score":0.021042943},"labels":[],"label_agreement":null},{"id":"W2071322747","doi":"10.1108/jd-02-2014-0037","title":"Differences over discourse structure differences: a reply to Urquhart and Urquhart","year":2015,"lang":"en","type":"article","venue":"Journal of Documentation","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Originality; Rhetorical question; Discipline; Epistemology; Strengths and weaknesses; Value (mathematics); Sociology; Computer science; Data science; Linguistics; Social science; Philosophy; Qualitative research","score_opus":0.03010225202968256,"score_gpt":0.31396739110756106,"score_spread":0.28386513907787847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071322747","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00065884955,0.005251387,0.0011963568,0.98569643,0.006185342,0.000007444077,0.000015274938,0.000019655128,0.0009693433],"genre_scores_gemma":[0.04823541,0.010458563,0.0035423604,0.91717094,0.016764753,0.00017741471,0.000034614666,0.00019849469,0.0034174093],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94132626,0.034866944,0.0034371712,0.007917305,0.011020604,0.0014318065],"domain_scores_gemma":[0.60149515,0.34325624,0.007219517,0.008540469,0.035555422,0.003933172],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07381124,0.0014368074,0.003163312,0.003586466,0.0112308385,0.013296708,0.006620803,0.03216164,0.004393784],"category_scores_gemma":[0.2542011,0.0011638752,0.0013902481,0.00609044,0.051263653,0.040640958,0.012857179,0.06401159,0.0023385922],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016754358,0.00008476666,0.00090273115,0.00059314276,0.00005660534,0.00033281595,0.06231893,0.00028284604,0.00021126117,0.2531515,0.6459357,0.035962027],"study_design_scores_gemma":[0.00013502513,0.00012251234,0.0014713397,0.003238317,0.00006643435,0.0007246703,0.048185483,0.001093774,0.0009288468,0.30150414,0.64218926,0.0003401706],"about_ca_topic_score_codex":0.011428738,"about_ca_topic_score_gemma":0.008140032,"teacher_disagreement_score":0.07381124,"about_ca_system_score_codex":0.010476593,"about_ca_system_score_gemma":0.011471176,"threshold_uncertainty_score":0.39035583},"labels":[],"label_agreement":null},{"id":"W2072473350","doi":"10.1109/icmla.2011.64","title":"Error Bounds for Online Predictions of Linear-Chain Conditional Random Fields: Application to Activity Recognition for Users of Rolling Walkers","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"CRFS; Conditional random field; Sequence (biology); Marginal distribution; Conditional probability distribution; Computer science; Class (philosophy); Observable; Conditional probability; Random variable; Distribution (mathematics); Chain (unit); Algorithm; Chain rule (probability); Artificial intelligence; Mathematics; Posterior probability; Regular conditional probability; Statistics; Mathematical analysis","score_opus":0.08022309942461758,"score_gpt":0.3007967272393659,"score_spread":0.22057362781474832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072473350","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028680302,0.00072096457,0.9673679,0.00048704646,0.0000719946,0.000040067738,0.00015986155,0.0014963457,0.00097562035],"genre_scores_gemma":[0.746391,0.0007851795,0.24762072,0.00025559633,0.00020796369,0.00022122107,0.0010367959,0.0006975035,0.0027840363],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968374,0.001133612,0.00019512807,0.00068993506,0.0008027955,0.00034111747],"domain_scores_gemma":[0.92058885,0.0676318,0.0025798674,0.004717452,0.0036452566,0.000836884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011224851,0.0015688227,0.0021202618,0.0011821712,0.0010259652,0.0020837411,0.0033501878,0.0020873286,0.0026927092],"category_scores_gemma":[0.06321311,0.00092328485,0.000893917,0.0012956903,0.002165755,0.004683429,0.002930887,0.004251881,0.00080696633],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031619103,0.000060420916,0.0016253667,0.00007232354,0.000036864327,0.00008498609,0.00010523405,0.95224047,0.0010622056,0.012216645,0.0010363511,0.031143015],"study_design_scores_gemma":[0.000002949224,0.000006825286,0.0000806393,0.0000057118627,0.0000020117272,0.000007747454,0.000005021509,0.99519014,0.0002945376,0.004334829,0.00006499719,0.0000045826564],"about_ca_topic_score_codex":0.01335648,"about_ca_topic_score_gemma":0.0098426975,"teacher_disagreement_score":0.01335648,"about_ca_system_score_codex":0.0027743278,"about_ca_system_score_gemma":0.0024010113,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2074335414","doi":"10.1145/2682571.2797088","title":"Efficient Computation of Co-occurrence Based Word Relatedness","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Word (group theory); Computer science; Pairwise comparison; Cluster analysis; Artificial intelligence; Natural language processing; Word lists by frequency; Computation; Word processing; Algorithm; Mathematics","score_opus":0.05896733454926951,"score_gpt":0.29747062204093633,"score_spread":0.23850328749166683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2074335414","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051803134,0.00042238046,0.93998903,0.00009960068,0.000045285793,0.00012018524,0.001308961,0.0041064424,0.0021050023],"genre_scores_gemma":[0.24559633,0.0003544307,0.7468849,0.000043446573,0.00007591343,0.00032425715,0.0046865703,0.00039636897,0.0016378039],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99821305,0.00038467546,0.00020780074,0.0004599747,0.0006405962,0.000093998155],"domain_scores_gemma":[0.99318576,0.0032609587,0.00066247996,0.0010960341,0.0016199059,0.00017489628],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011535326,0.00063456286,0.0012077603,0.0066206404,0.000731904,0.0017782451,0.00095663476,0.0005482644,0.0023803057],"category_scores_gemma":[0.01719649,0.00044305774,0.0006200082,0.00607259,0.00042168322,0.0032894586,0.0014669537,0.0008439755,0.002910065],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027285112,0.0001999716,0.011172283,0.00044970386,0.00019294146,0.00022623388,0.00084920967,0.029324904,0.046808325,0.018193616,0.007786832,0.88452315],"study_design_scores_gemma":[0.000057234887,0.0003380419,0.021828553,0.00009157079,0.00016884407,0.0011153936,0.0008502333,0.78729296,0.05317496,0.11048242,0.024436876,0.00016283475],"about_ca_topic_score_codex":0.0019014394,"about_ca_topic_score_gemma":0.0029107996,"teacher_disagreement_score":0.0066206404,"about_ca_system_score_codex":0.0005131591,"about_ca_system_score_gemma":0.0013811897,"threshold_uncertainty_score":0.007962942},"labels":[],"label_agreement":null},{"id":"W2074811005","doi":"10.1145/2633211.2634351","title":"Tulip","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Centroid; Information retrieval; Feature (linguistics); Entity linking; Set (abstract data type); Context (archaeology); Feature vector; Graph; Artificial intelligence; Natural language processing; Theoretical computer science; Knowledge base; Linguistics; Programming language","score_opus":0.01732593325240539,"score_gpt":0.221161307912772,"score_spread":0.2038353746603666,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2074811005","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0106690535,0.00611884,0.12787008,0.003413891,0.0032269494,0.000856394,0.095673315,0.4677415,0.28442997],"genre_scores_gemma":[0.05857527,0.004417366,0.16447917,0.003937368,0.00089118903,0.0010553582,0.44154674,0.050775725,0.27432185],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978187,0.00025517296,0.00018956172,0.00057868747,0.0008899322,0.00026783856],"domain_scores_gemma":[0.99732494,0.00042528726,0.00010380867,0.0009304471,0.0009769761,0.0002386419],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0018796059,0.0018831894,0.0014186003,0.003745595,0.0018526404,0.004875112,0.0032716403,0.0020548904,0.1150427],"category_scores_gemma":[0.006023239,0.0008915667,0.0013898629,0.0026230498,0.00048078803,0.0063957646,0.0061951797,0.0021562048,0.20641857],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068931206,0.00012802459,0.00072623865,0.00076231884,0.000058212394,0.0002732077,0.00022749988,0.0009081902,0.004676017,0.0065446063,0.73013884,0.25486758],"study_design_scores_gemma":[0.00007429067,0.00010627317,0.00089320284,0.00013092285,0.00004707737,0.00046801378,0.000111416346,0.0078682965,0.0063245613,0.0061424337,0.97774225,0.000091337264],"about_ca_topic_score_codex":0.0048596906,"about_ca_topic_score_gemma":0.0050667822,"teacher_disagreement_score":0.8849573,"about_ca_system_score_codex":0.0013648578,"about_ca_system_score_gemma":0.0016356481,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2075245318","doi":"10.1145/1670564.1670576","title":"TREC-CHEM","year":2009,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Scalability; Data science; Domain (mathematical analysis); Information retrieval; Database","score_opus":0.019420468175832994,"score_gpt":0.25382635029268097,"score_spread":0.23440588211684799,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2075245318","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00882177,0.009882298,0.010506333,0.013251288,0.013620745,0.0033541112,0.7666656,0.016390178,0.15750769],"genre_scores_gemma":[0.010286338,0.0021103185,0.013331952,0.002667142,0.001247173,0.0014180993,0.83546954,0.001988428,0.13148095],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99419963,0.0014952304,0.00045446825,0.0007230131,0.0025398366,0.0005877276],"domain_scores_gemma":[0.9796708,0.002471416,0.0007761234,0.0021890493,0.0126862805,0.0022062222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009740395,0.0029343741,0.0029612854,0.0069714,0.0023002601,0.0048875134,0.0045259283,0.0034914033,0.15126272],"category_scores_gemma":[0.015590087,0.00078370684,0.0016185389,0.0045351735,0.00092661876,0.0042739245,0.0024653818,0.0035044176,0.122069314],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016200835,0.000079095764,0.00020769397,0.0002667805,0.000033766937,0.000022189544,0.0000138218375,0.000447418,0.00062722573,0.00058477785,0.9836056,0.013949608],"study_design_scores_gemma":[0.000573419,0.00029392936,0.004657444,0.00027578525,0.00013407171,0.00018846475,0.000114476774,0.003402325,0.0050379154,0.0036748003,0.9815112,0.00013626982],"about_ca_topic_score_codex":0.06540788,"about_ca_topic_score_gemma":0.098365664,"teacher_disagreement_score":0.15126272,"about_ca_system_score_codex":0.00470002,"about_ca_system_score_gemma":0.00950511,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2076039115","doi":"10.1109/tkde.2009.25","title":"Evaluating the Generation of Domain Ontologies in the Knowledge Puzzle Project","year":2009,"lang":"en","type":"article","venue":"IEEE Transactions on Knowledge and Data Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Ontology; Computer science; Ontology learning; Upper ontology; Information retrieval; Domain (mathematical analysis); Ontology-based data integration; Domain knowledge; Process ontology; Suggested Upper Merged Ontology; Set (abstract data type); Ontology alignment; Natural language processing; Data science; Artificial intelligence; Semantic Web","score_opus":0.1546157879818343,"score_gpt":0.36613754980686486,"score_spread":0.21152176182503055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2076039115","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8988647,0.000856367,0.07929979,0.0006158837,0.00012778924,0.0013426145,0.002410328,0.0046237865,0.011858734],"genre_scores_gemma":[0.67460537,0.0005144653,0.30424497,0.00022025654,0.00003265066,0.001032556,0.014230281,0.0006784412,0.0044409814],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98950785,0.005993048,0.0008969897,0.001039462,0.002334165,0.00022856523],"domain_scores_gemma":[0.941774,0.04761125,0.0015121897,0.003718784,0.004642796,0.0007411059],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011352773,0.0010837074,0.00089521427,0.0025746713,0.0008788874,0.0023253309,0.001997393,0.001963608,0.0028791174],"category_scores_gemma":[0.06540478,0.0004931436,0.00058970996,0.002676472,0.0011830695,0.004852819,0.0035030225,0.0012406069,0.00075598276],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043143257,0.0038382378,0.018344767,0.0031201944,0.00043628964,0.0016273275,0.006198972,0.16609229,0.018576037,0.01264454,0.022782387,0.74202466],"study_design_scores_gemma":[0.0014745524,0.0036329068,0.026520375,0.00040295455,0.00031869914,0.0013029085,0.006953436,0.81965816,0.07659437,0.013284235,0.049661987,0.00019542685],"about_ca_topic_score_codex":0.005153508,"about_ca_topic_score_gemma":0.0051202616,"teacher_disagreement_score":0.011352773,"about_ca_system_score_codex":0.0018253069,"about_ca_system_score_gemma":0.0016431169,"threshold_uncertainty_score":0.060039937},"labels":[],"label_agreement":null},{"id":"W2077054502","doi":"10.1145/2396761.2398676","title":"Improving the performance of the reinforcement learning model for answering complex questions","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Reinforcement learning; Automatic summarization; Computer science; Question answering; Artificial intelligence; Task (project management); Set (abstract data type); Sentence; Natural language processing; Process (computing); Reinforcement; Feature (linguistics); Machine learning","score_opus":0.04739034023367434,"score_gpt":0.2605318500398405,"score_spread":0.21314150980616614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2077054502","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16664651,0.00083616783,0.8254575,0.0006952152,0.000094897994,0.000115725954,0.00009756299,0.0037495107,0.002307042],"genre_scores_gemma":[0.86879396,0.00018350649,0.1283477,0.00019213825,0.000083736144,0.00009640784,0.0003013448,0.00012858138,0.0018726087],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986541,0.00068884646,0.00006279467,0.0002946906,0.0002070469,0.00009257467],"domain_scores_gemma":[0.9922335,0.005844942,0.00031718172,0.0005270423,0.00085845054,0.00021875702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042709345,0.0009127955,0.0013601001,0.00044079323,0.00038054204,0.0008793787,0.001515688,0.0014422573,0.0013472505],"category_scores_gemma":[0.016414946,0.00032148894,0.00045698037,0.0003590105,0.00040684117,0.00225204,0.0009810845,0.0019080698,0.0007170714],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007910152,0.00067482685,0.0041403943,0.0001527205,0.00010344245,0.00013157097,0.00030193033,0.68686175,0.011771882,0.003910266,0.003615915,0.28754434],"study_design_scores_gemma":[0.000012577936,0.000047738937,0.000119404205,0.0000013553529,0.0000038342105,0.000006093207,0.0000050188078,0.9981375,0.0007433538,0.0008029216,0.00011686736,0.000003449475],"about_ca_topic_score_codex":0.0061266352,"about_ca_topic_score_gemma":0.0041427584,"teacher_disagreement_score":0.0061266352,"about_ca_system_score_codex":0.00084770756,"about_ca_system_score_gemma":0.0010405509,"threshold_uncertainty_score":0.02258718},"labels":[],"label_agreement":null},{"id":"W2078923208","doi":"10.1145/2494266.2494297","title":"Beyond term clusters","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Information retrieval; Topic model; Digital library; Domain (mathematical analysis); Term (time); Range (aeronautics); Data science; Mathematics","score_opus":0.014061333726225308,"score_gpt":0.22235147737319957,"score_spread":0.20829014364697426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2078923208","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0112795895,0.0021751204,0.9705508,0.0016801811,0.000476878,0.00026633858,0.0021015978,0.002470949,0.008998556],"genre_scores_gemma":[0.29400837,0.0027670222,0.6566773,0.001183697,0.00146476,0.001179721,0.008964006,0.0019087729,0.0318464],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958702,0.0010060319,0.00022582816,0.00150162,0.0011096142,0.0002866931],"domain_scores_gemma":[0.98995054,0.0045732316,0.00066345057,0.0024175572,0.0018963467,0.00049889303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039016528,0.001622448,0.001951055,0.005770852,0.002884919,0.005892384,0.00550318,0.0029531394,0.0084964195],"category_scores_gemma":[0.023008846,0.0010229192,0.0023860773,0.007822914,0.0019953866,0.011852444,0.0043487037,0.0033808805,0.0062326496],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006071575,0.00026415303,0.0068945014,0.0008473244,0.0005165231,0.0003085165,0.0018643882,0.10521952,0.0030702648,0.39934927,0.066700414,0.41435796],"study_design_scores_gemma":[0.000049664246,0.00006490436,0.0011038904,0.00011530589,0.00013379686,0.00025619034,0.00024594297,0.427935,0.0016168212,0.5219283,0.04647408,0.00007604069],"about_ca_topic_score_codex":0.011497587,"about_ca_topic_score_gemma":0.013681786,"teacher_disagreement_score":0.011497587,"about_ca_system_score_codex":0.0024098458,"about_ca_system_score_gemma":0.0032568472,"threshold_uncertainty_score":0.028423429},"labels":[],"label_agreement":null},{"id":"W2079629183","doi":"10.3115/1609067.1609136","title":"Using lexical and relational similarity to classify semantic relations","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Engineering and Physical Sciences Research Council; Atomic Energy of Canada Limited","keywords":"Computer science; Natural language processing; Semantic similarity; Artificial intelligence; Similarity (geometry); Task (project management); Relational database; Word (group theory); Noun; Information retrieval; Linguistics","score_opus":0.10996154736035278,"score_gpt":0.3180892979233171,"score_spread":0.2081277505629643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079629183","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060544554,0.0006678718,0.9327547,0.0002890248,0.000048008318,0.00022492559,0.0005447586,0.0010247045,0.0039015377],"genre_scores_gemma":[0.71610624,0.0004569323,0.2795064,0.000086843254,0.0001038583,0.0002446215,0.0016681547,0.0001269117,0.0017001],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995934,0.000962176,0.00049783255,0.0010475169,0.0013457126,0.00021269353],"domain_scores_gemma":[0.9946043,0.0023832351,0.00069007545,0.0013454068,0.0007881773,0.0001887841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028344241,0.000789846,0.001132261,0.008319267,0.0009134162,0.0035533651,0.0018694423,0.0013710656,0.002605613],"category_scores_gemma":[0.013556665,0.0003764435,0.0017008074,0.0053302003,0.001750606,0.009945053,0.0031424055,0.0011512581,0.0010015683],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006151776,0.00061314524,0.029244334,0.0006682347,0.000693717,0.00041653335,0.0025758115,0.058399227,0.021035735,0.256658,0.0049100337,0.62417006],"study_design_scores_gemma":[0.000052639152,0.00021356584,0.006883624,0.000083367544,0.0002497126,0.00057627296,0.0006555205,0.7342036,0.0074041397,0.2427954,0.006762416,0.00011973117],"about_ca_topic_score_codex":0.0040805303,"about_ca_topic_score_gemma":0.003768182,"teacher_disagreement_score":0.008319267,"about_ca_system_score_codex":0.0011257249,"about_ca_system_score_gemma":0.000999836,"threshold_uncertainty_score":0.014990032},"labels":[],"label_agreement":null},{"id":"W2080378089","doi":"10.3758/bf03196761","title":"Using context to build semantics","year":2005,"lang":"en","type":"article","venue":"Psychonomic Bulletin & Review","topic":"Topic Modeling","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada","funders":"","keywords":"Latent semantic analysis; Semantics (computer science); Dimension (graph theory); Context (archaeology); Representation (politics); Psychology; Similarity (geometry); Natural language processing; Semantic memory; Singular value decomposition; Semantic similarity; Episodic memory; Artificial intelligence; Computer science; Cognition; Mathematics","score_opus":0.059725783596568195,"score_gpt":0.32477373437422885,"score_spread":0.26504795077766063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080378089","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019512873,0.0047562835,0.95889544,0.0021512448,0.00051816826,0.00017320877,0.0022727805,0.0039130948,0.007806968],"genre_scores_gemma":[0.35746422,0.0033882614,0.62920326,0.0005633891,0.00037617865,0.00032384507,0.0058150785,0.0008001504,0.0020656567],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9976828,0.0009709797,0.00025620396,0.0006546923,0.00034110644,0.00009424974],"domain_scores_gemma":[0.9955309,0.0025894505,0.0003035339,0.0008773973,0.0005317841,0.00016688227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029240614,0.0011834968,0.0010274017,0.0063183843,0.0014954295,0.004620035,0.0015031506,0.0010207342,0.004191911],"category_scores_gemma":[0.011882886,0.00093292125,0.0029420783,0.0032576027,0.0012913838,0.011082153,0.0029577212,0.0027722428,0.0015301701],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041470397,0.0001370647,0.005030194,0.0019996779,0.00078319194,0.0004741105,0.002963777,0.017545583,0.0061716307,0.55716604,0.017154528,0.39015952],"study_design_scores_gemma":[0.000061175,0.00005439341,0.0010417304,0.0003823157,0.0003798299,0.0002014618,0.0003766897,0.05729145,0.0035339969,0.8727804,0.063828625,0.00006798529],"about_ca_topic_score_codex":0.003132469,"about_ca_topic_score_gemma":0.005541226,"teacher_disagreement_score":0.0063183843,"about_ca_system_score_codex":0.0013084774,"about_ca_system_score_gemma":0.002114021,"threshold_uncertainty_score":0.015464067},"labels":[],"label_agreement":null},{"id":"W2081877265","doi":"10.1007/s10115-012-0504-y","title":"An efficient concept-based retrieval model for enhancing text retrieval quality","year":2012,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; D2L (Canada)","funders":"","keywords":"Computer science; Sentence; Information retrieval; Natural language processing; Weighting; Term (time); Semantics (computer science); Latent semantic analysis; Term Discrimination; Representation (politics); Artificial intelligence; Graph; Document retrieval; Relevance (law); Concept search; Theoretical computer science; Search engine","score_opus":0.03897948914390884,"score_gpt":0.30871514418177404,"score_spread":0.2697356550378652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081877265","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019612933,0.00050500006,0.977948,0.00024109616,0.000063803025,0.00009121676,0.00017481153,0.0007645087,0.0005987367],"genre_scores_gemma":[0.47244087,0.00083957583,0.52089435,0.00025927986,0.00025051567,0.000491984,0.0008192457,0.00020006057,0.0038041463],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99901664,0.0002965868,0.00008754757,0.0002118111,0.00031230357,0.00007519556],"domain_scores_gemma":[0.9981437,0.00084968505,0.00012005099,0.00017684461,0.00065659825,0.000053110074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020468875,0.00066446216,0.0013112675,0.0015422367,0.00048855704,0.0014082346,0.0015493894,0.0012832701,0.0018776093],"category_scores_gemma":[0.0050486275,0.00036478005,0.0010325984,0.0015833902,0.0004336629,0.0035721618,0.00073849614,0.0011017387,0.0010499181],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010851385,0.0006749158,0.0020586823,0.0005953547,0.00028738604,0.0001487729,0.00032101315,0.3070586,0.05074766,0.0392626,0.0116788335,0.5860811],"study_design_scores_gemma":[0.000031218948,0.00006080605,0.00025680193,0.0000069099597,0.000054412292,0.0000458482,0.000011072073,0.99156123,0.0030079228,0.004221866,0.0007254585,0.000016501228],"about_ca_topic_score_codex":0.0042777737,"about_ca_topic_score_gemma":0.004198247,"teacher_disagreement_score":0.0042777737,"about_ca_system_score_codex":0.0010475179,"about_ca_system_score_gemma":0.001399417,"threshold_uncertainty_score":0.010825098},"labels":[],"label_agreement":null},{"id":"W20831368","doi":"10.3233/978-1-60750-028-5-49","title":"Predicting Learner Answers Correctness through Brainwaves Assesment and Emotional Dimensions","year":2009,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Correctness; Psychology; Computer science; Algorithm","score_opus":0.04530497537496673,"score_gpt":0.2819834301727603,"score_spread":0.23667845479779356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W20831368","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96223724,0.000327407,0.03352164,0.0001255378,0.000030953208,0.00007906538,0.0006362948,0.00033665381,0.0027051682],"genre_scores_gemma":[0.97612983,0.00020142065,0.019847,0.000027086655,0.000022570533,0.00006929285,0.00090777193,0.000031763815,0.0027632937],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99929667,0.00028612133,0.000057871144,0.00013724231,0.00017627086,0.00004572726],"domain_scores_gemma":[0.9908522,0.0072999555,0.0005904512,0.00021714845,0.0008114617,0.00022881027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014984338,0.000480148,0.00035462595,0.0009608311,0.000114762646,0.0008036842,0.0002992823,0.0006349186,0.0028308223],"category_scores_gemma":[0.0115578445,0.00011556396,0.000302909,0.0004883168,0.00014495882,0.0008394159,0.0003499193,0.00044723955,0.0015257441],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071608793,0.00043393223,0.67573,0.00024155434,0.00012402289,0.00020332915,0.0007172119,0.00721821,0.015767317,0.00030805432,0.0021195437,0.29642078],"study_design_scores_gemma":[0.00004558968,0.0013310607,0.7544175,0.00008097546,0.00014293406,0.0006073223,0.0009186312,0.21361847,0.02425094,0.0014024588,0.003101921,0.00008217302],"about_ca_topic_score_codex":0.00080227206,"about_ca_topic_score_gemma":0.0010135039,"teacher_disagreement_score":0.0028308223,"about_ca_system_score_codex":0.000160313,"about_ca_system_score_gemma":0.000146557,"threshold_uncertainty_score":0.009470105},"labels":[],"label_agreement":null},{"id":"W2083365005","doi":"10.1371/journal.pone.0025085","title":"Judgment of the Humanness of an Interlocutor Is in the Eye of the Beholder","year":2011,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"Natural Sciences and Engineering Research Council of Canada; Université Laval","keywords":"Psychology; Cognition; Social psychology; Cognitive psychology","score_opus":0.13557126643027673,"score_gpt":0.24047733562551474,"score_spread":0.10490606919523801,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2083365005","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9673867,0.00015900347,0.023455648,0.0003207416,0.000019354962,0.00002004912,0.000031966985,0.0000530214,0.008553378],"genre_scores_gemma":[0.9970023,0.00003200437,0.002439892,0.00003583704,0.0000063329776,0.000005619564,0.000015452955,0.000012550434,0.00044997543],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9972868,0.0017334722,0.00007843904,0.00037659283,0.00041146373,0.00011316311],"domain_scores_gemma":[0.9879682,0.0071193627,0.002240867,0.0011829915,0.0010753192,0.0004132781],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025073802,0.00027254564,0.000335443,0.00067621976,0.0008585259,0.0023374443,0.00032633246,0.0006170639,0.001723165],"category_scores_gemma":[0.014590879,0.00020313778,0.00020442445,0.00035042406,0.0024226233,0.0016630042,0.00106875,0.00082712015,0.00028215133],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001991422,0.00029366204,0.2829594,0.00060996955,0.00034187408,0.0013255929,0.18888868,0.00735854,0.32243493,0.06969081,0.0027762987,0.121328786],"study_design_scores_gemma":[0.000053358857,0.00090372004,0.6439648,0.00020614221,0.0002774743,0.0016491164,0.11350714,0.06350018,0.069101416,0.08917558,0.017232375,0.00042859983],"about_ca_topic_score_codex":0.0017007225,"about_ca_topic_score_gemma":0.0018133076,"teacher_disagreement_score":0.0025073802,"about_ca_system_score_codex":0.00044537036,"about_ca_system_score_gemma":0.00038575733,"threshold_uncertainty_score":0.013260424},"labels":[],"label_agreement":null},{"id":"W2084105113","doi":"10.1016/s1532-0464(03)00016-9","title":"Paraphrasing for condensation in journal abstracting","year":2002,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Paraphrase; Computer science; Natural language processing; Generality; Rhetorical question; Artificial intelligence; Variety (cybernetics); Linguistics; Sentence; Sublanguage; Psychology","score_opus":0.051266578281694015,"score_gpt":0.28873543523871736,"score_spread":0.23746885695702336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084105113","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04494936,0.0024349669,0.904123,0.003762477,0.0017651486,0.0007327583,0.004668366,0.008375756,0.029188208],"genre_scores_gemma":[0.56289816,0.0011877476,0.41176188,0.000796354,0.0015889115,0.0005996383,0.007540764,0.002613718,0.01101284],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99104404,0.0047981786,0.0009777326,0.001138595,0.0016744932,0.0003668673],"domain_scores_gemma":[0.96581507,0.021092592,0.0017804081,0.0046956534,0.006007617,0.0006087256],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0048784073,0.0015786098,0.00094666745,0.0058635813,0.0024422086,0.004005141,0.0018673029,0.0015722368,0.022824906],"category_scores_gemma":[0.044055667,0.00094108825,0.00132003,0.005798946,0.0017902205,0.008594322,0.0045720125,0.0021878134,0.00582924],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013179209,0.00021724761,0.0032669213,0.0026684096,0.0002436276,0.0014065171,0.016336795,0.0052208337,0.03933402,0.2274037,0.13049024,0.5720938],"study_design_scores_gemma":[0.00035953237,0.00041756153,0.005488214,0.00086127076,0.00051701214,0.0019639465,0.0061992104,0.14856705,0.04833935,0.37370145,0.41328982,0.00029556814],"about_ca_topic_score_codex":0.0019945,"about_ca_topic_score_gemma":0.0021955476,"teacher_disagreement_score":0.9951216,"about_ca_system_score_codex":0.0015620741,"about_ca_system_score_gemma":0.0017231356,"threshold_uncertainty_score":0.07635695},"labels":[],"label_agreement":null},{"id":"W2086038979","doi":"10.1016/j.eswa.2014.09.015","title":"A novel contextual topic model for multi-document summarization","year":2014,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Information overload; Multi-document summarization; Context (archaeology); Artificial intelligence; Word (group theory); Topic model; Natural language processing; Information retrieval; World Wide Web; Linguistics","score_opus":0.0467911233113684,"score_gpt":0.290877761885536,"score_spread":0.2440866385741676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086038979","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0052541588,0.0012772891,0.990316,0.00024334562,0.00017521712,0.00008445759,0.000512331,0.0014893711,0.0006477997],"genre_scores_gemma":[0.25603142,0.0029221221,0.7240017,0.00044884757,0.0014983652,0.0010126295,0.006421034,0.00086461415,0.006799313],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981667,0.0006163746,0.00018945521,0.00052213663,0.000355089,0.0001502247],"domain_scores_gemma":[0.99769884,0.0011544317,0.00015767732,0.000269049,0.0006256465,0.000094324096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021575296,0.0012695198,0.0018491503,0.002159778,0.0008566994,0.0020300043,0.0019183551,0.0017958892,0.0027212184],"category_scores_gemma":[0.0058393655,0.0006406002,0.0019281922,0.0030874969,0.00041076043,0.0031439385,0.0013883811,0.0019738034,0.0026854505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090526976,0.0003247783,0.0027442712,0.0010499435,0.0006306788,0.00030466542,0.0007537971,0.14267734,0.027241476,0.028503949,0.028235562,0.76662827],"study_design_scores_gemma":[0.000039331793,0.00010343354,0.0007251717,0.000038017228,0.00020340494,0.00011920987,0.00007247814,0.97542113,0.0035113413,0.01113267,0.008588348,0.000045511428],"about_ca_topic_score_codex":0.005175577,"about_ca_topic_score_gemma":0.008498368,"teacher_disagreement_score":0.005175577,"about_ca_system_score_codex":0.00071975077,"about_ca_system_score_gemma":0.0014328198,"threshold_uncertainty_score":0.011410236},"labels":[],"label_agreement":null},{"id":"W2088772104","doi":"10.1145/2063576.2063879","title":"Recommending citations with translation model","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Pratt and Whitney Canada","keywords":"Computer science; Citation; Bridge (graph theory); Translation (biology); Vocabulary; Natural language processing; Information retrieval; Artificial intelligence; World Wide Web; Linguistics","score_opus":0.1769830511217303,"score_gpt":0.2633044284310369,"score_spread":0.08632137730930661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088772104","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13557968,0.00727838,0.83297807,0.0041450514,0.0011633064,0.0006674951,0.0021308388,0.0044205785,0.0116366595],"genre_scores_gemma":[0.7390042,0.0038531132,0.2246473,0.0011225085,0.0020279167,0.0010127228,0.0052813534,0.0004621071,0.022588855],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99706894,0.0012476351,0.00024139803,0.00077701535,0.00048941054,0.00017557574],"domain_scores_gemma":[0.9928571,0.0046933517,0.00036901282,0.000582578,0.001261288,0.0002365954],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0030801718,0.0015237305,0.0026955814,0.0071858773,0.0016848528,0.0026373635,0.0020018504,0.0033948566,0.005997873],"category_scores_gemma":[0.016886733,0.0006429122,0.0022957167,0.009665354,0.0007278236,0.0045298547,0.0011581172,0.0017429704,0.0031536722],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009715409,0.0008845918,0.014346331,0.0012401667,0.00072452857,0.00067073345,0.000699152,0.1769468,0.0049219783,0.033478245,0.0525193,0.71259665],"study_design_scores_gemma":[0.00013400262,0.00013125509,0.0012584246,0.000039111474,0.00020120914,0.00020701476,0.00007179597,0.96869266,0.0014932972,0.022324145,0.0053928513,0.00005434235],"about_ca_topic_score_codex":0.00712329,"about_ca_topic_score_gemma":0.0068686362,"teacher_disagreement_score":0.9928141,"about_ca_system_score_codex":0.0010862739,"about_ca_system_score_gemma":0.0020278455,"threshold_uncertainty_score":0.02006489},"labels":[],"label_agreement":null},{"id":"W2091812280","doi":"10.1145/1273496.1273577","title":"Three new graphical models for statistical language modelling","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":582,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Probabilistic logic; Language model; Representation (politics); Word (group theory); n-gram; Binary number; Set (abstract data type); Parametric statistics; Artificial intelligence; Natural language processing; Statistical model; Sequence (biology); Theoretical computer science; Probability distribution; Mathematics; Programming language","score_opus":0.05588757867207705,"score_gpt":0.29496802726042914,"score_spread":0.2390804485883521,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091812280","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018124898,0.00035550186,0.9948315,0.0008353207,0.00007483315,0.000038858634,0.0003021453,0.0005263443,0.001222971],"genre_scores_gemma":[0.2039946,0.0019260241,0.78098744,0.0012635596,0.0005781467,0.0007144646,0.0016929578,0.0004518624,0.008390965],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968266,0.0014224574,0.00021870167,0.0005622492,0.00077634805,0.00019367838],"domain_scores_gemma":[0.9935021,0.0042546936,0.00047067148,0.00089493056,0.00062413816,0.0002534551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037155575,0.0013534448,0.0012556722,0.0021964323,0.00062799884,0.0030535213,0.003442994,0.0024171073,0.0045342674],"category_scores_gemma":[0.016416732,0.00079775933,0.0026416213,0.0022535236,0.0017101136,0.0046954053,0.0033791808,0.0033080836,0.0019421356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018139266,0.00010160913,0.0015067401,0.0002694862,0.00018524464,0.0002782455,0.0004595977,0.15296814,0.002838111,0.6932817,0.009388517,0.13854133],"study_design_scores_gemma":[0.000034261833,0.00002872049,0.00021140427,0.000029844085,0.000043627777,0.00013327456,0.000027027685,0.59997535,0.0004526685,0.39152643,0.00749618,0.00004134311],"about_ca_topic_score_codex":0.004173832,"about_ca_topic_score_gemma":0.0056059603,"teacher_disagreement_score":0.0045342674,"about_ca_system_score_codex":0.001575413,"about_ca_system_score_gemma":0.0012189589,"threshold_uncertainty_score":0.019649982},"labels":[],"label_agreement":null},{"id":"W2092057061","doi":"10.3758/bf03193815","title":"Similes on the Internet have explanations","year":2006,"lang":"en","type":"article","venue":"Psychonomic Bulletin & Review","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"","keywords":"Psychology; The Internet; Disease; Relation (database); Social psychology; World Wide Web; Medicine; Pathology; Computer science","score_opus":0.033352978128346186,"score_gpt":0.262731170050443,"score_spread":0.2293781919220968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2092057061","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.102811575,0.10370642,0.10284282,0.114161335,0.009587709,0.00023726062,0.0025341045,0.001338395,0.5627803],"genre_scores_gemma":[0.8661472,0.036372,0.017878426,0.0121167395,0.008793982,0.00019577911,0.0019108091,0.00037690258,0.05620827],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.99939585,0.0002593726,0.0000414015,0.00014341016,0.000098938806,0.00006099699],"domain_scores_gemma":[0.9942778,0.003725765,0.000585561,0.00082707766,0.0004017392,0.0001822287],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008625992,0.000709668,0.00054722937,0.0031843688,0.0010955642,0.0030281127,0.00089335965,0.002392706,0.020599077],"category_scores_gemma":[0.007012048,0.00023670294,0.0006068167,0.0035505837,0.0031447548,0.008930414,0.0014217523,0.0021824848,0.0022201352],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013737414,0.00005251032,0.00470203,0.0011485032,0.00011828576,0.00069254515,0.004905594,0.00061810855,0.0006417552,0.845994,0.0583196,0.08266974],"study_design_scores_gemma":[0.0000871238,0.000073381765,0.012997473,0.000701501,0.00013326273,0.0021397816,0.0037315928,0.0027752377,0.0010303415,0.54933685,0.4269356,0.00005774039],"about_ca_topic_score_codex":0.0012241884,"about_ca_topic_score_gemma":0.0014502738,"teacher_disagreement_score":0.020599077,"about_ca_system_score_codex":0.0006622858,"about_ca_system_score_gemma":0.0002509747,"threshold_uncertainty_score":0.06891072},"labels":[],"label_agreement":null},{"id":"W2092922846","doi":"10.1145/2344416.2344418","title":"Extracting information networks from the blogosphere","year":2012,"lang":"en","type":"article","venue":"ACM Transactions on the Web","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Benchmark (surveying); Relationship extraction; Task (project management); Information extraction; Pruning; Blogosphere; Weighting; Cluster analysis; Data mining; Information retrieval; Relation (database); Filter (signal processing); Artificial intelligence; The Internet; World Wide Web","score_opus":0.025162967948974428,"score_gpt":0.2268239240225655,"score_spread":0.20166095607359105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2092922846","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38507882,0.008543748,0.5388272,0.0017916244,0.0003834404,0.0004836392,0.03487159,0.007252821,0.022767145],"genre_scores_gemma":[0.56285053,0.0058565377,0.37305757,0.00018274604,0.00048831414,0.00039121162,0.050699785,0.000495005,0.0059783277],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993223,0.00012929353,0.00006362869,0.0001617475,0.00023266976,0.000090418565],"domain_scores_gemma":[0.99679095,0.0019087619,0.00043709786,0.000331485,0.00042174707,0.000109990295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070880353,0.0011539254,0.0006520719,0.0109303715,0.00074612995,0.0018064711,0.0005914315,0.0007644632,0.0012867367],"category_scores_gemma":[0.0058668484,0.00040114895,0.0006285653,0.00897022,0.0003814068,0.003812918,0.00121469,0.0005540692,0.0013878054],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045343244,0.00031721653,0.04104189,0.0022474197,0.00024075668,0.0018247275,0.001459956,0.042553376,0.04130438,0.02142651,0.03612445,0.811006],"study_design_scores_gemma":[0.00007332953,0.00024833548,0.06793076,0.0005143509,0.00036655625,0.002457115,0.0027364555,0.5438445,0.0662845,0.12379277,0.19158241,0.00016899788],"about_ca_topic_score_codex":0.0033283134,"about_ca_topic_score_gemma":0.006778314,"teacher_disagreement_score":0.0109303715,"about_ca_system_score_codex":0.0005724587,"about_ca_system_score_gemma":0.00079699873,"threshold_uncertainty_score":0.0066179037},"labels":[],"label_agreement":null},{"id":"W2093641143","doi":"10.3115/1073445.1073477","title":"Frequency estimates for statistical word similarity measures","year":2003,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":197,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Coordenação de Aperfeiçoamento de Pessoal de Nível Superior","keywords":"Computer science; Word (group theory); Word lists by frequency; Natural language processing; Similarity (geometry); Artificial intelligence; Context (archaeology); Synonym (taxonomy); Semantic similarity; Set (abstract data type); Word2vec; Information retrieval; Sentence; Linguistics","score_opus":0.058525926575785155,"score_gpt":0.29914493728641595,"score_spread":0.2406190107106308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2093641143","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03205218,0.0011203255,0.96341836,0.00012378229,0.000074205476,0.00021473212,0.00073527556,0.0008311162,0.0014299721],"genre_scores_gemma":[0.31321496,0.00077715836,0.6808929,0.000099790814,0.0003702739,0.001072166,0.0024461756,0.00034045582,0.00078605197],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9829086,0.0073927585,0.0015412305,0.0021331618,0.005728808,0.00029544297],"domain_scores_gemma":[0.8117186,0.15942875,0.007500245,0.012592444,0.008221723,0.0005383765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01439873,0.0010780077,0.0012861289,0.0130980015,0.0007595656,0.0025936512,0.0017473857,0.001716024,0.0033769414],"category_scores_gemma":[0.1654783,0.00063672353,0.0010301356,0.007947293,0.0012112533,0.0068152766,0.0017346035,0.0017843456,0.0015091103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072171417,0.00043734204,0.036699727,0.0015181267,0.0007909104,0.00023253124,0.001331701,0.111428976,0.014985406,0.11268343,0.0070077833,0.71216244],"study_design_scores_gemma":[0.00018565918,0.0006005115,0.027598316,0.00033305484,0.00018609756,0.0013937884,0.00074000325,0.7638283,0.01956507,0.17075855,0.014490938,0.00031979667],"about_ca_topic_score_codex":0.0010883367,"about_ca_topic_score_gemma":0.00090390054,"teacher_disagreement_score":0.01439873,"about_ca_system_score_codex":0.0010480119,"about_ca_system_score_gemma":0.00074320444,"threshold_uncertainty_score":0.07614863},"labels":[],"label_agreement":null},{"id":"W2094723464","doi":"10.1109/nlpke.2007.4368014","title":"Recognizing Biomedical Named Entities in the Absence of Human Annotated Corpora","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Annotation; Artificial intelligence; Classifier (UML); Task (project management); Named-entity recognition; Natural language processing; Training set; Support vector machine; Domain (mathematical analysis); Process (computing); Supervised learning; Labeled data; Machine learning; Information retrieval","score_opus":0.046766609992947863,"score_gpt":0.2905173275352811,"score_spread":0.24375071754233324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2094723464","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13505358,0.0037427829,0.83634293,0.002684401,0.00045447386,0.00035402362,0.0072212173,0.007253833,0.006892732],"genre_scores_gemma":[0.32721698,0.0015122972,0.6300633,0.00064368104,0.000488594,0.0006322581,0.03485103,0.0005501114,0.004041754],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9920642,0.0039260997,0.0006827052,0.0019970846,0.0011148986,0.00021506358],"domain_scores_gemma":[0.95073104,0.0364552,0.0031740149,0.006297452,0.0028429392,0.0004993912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010074201,0.0012415757,0.0017876882,0.0041067125,0.001239126,0.0023937423,0.0023961721,0.0028638833,0.0015630189],"category_scores_gemma":[0.035185162,0.0007488903,0.0008610955,0.0035195164,0.0014178954,0.006919385,0.002612598,0.001852123,0.0022911013],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012787615,0.0007320777,0.02317154,0.0027169304,0.00037606078,0.0031964849,0.0025302838,0.04802118,0.0837971,0.016839337,0.05188429,0.7654559],"study_design_scores_gemma":[0.00016050308,0.0004026663,0.018945394,0.00043605096,0.00033518104,0.003507146,0.002180866,0.69562733,0.10991752,0.0575114,0.11078611,0.00018983604],"about_ca_topic_score_codex":0.0014936464,"about_ca_topic_score_gemma":0.0037655819,"teacher_disagreement_score":0.010074201,"about_ca_system_score_codex":0.0007551591,"about_ca_system_score_gemma":0.0014614315,"threshold_uncertainty_score":0.05327809},"labels":[],"label_agreement":null},{"id":"W2096537984","doi":"","title":"Using coreference links and sentence compression in graph-based summarization","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Coreference; Computer science; Sentence; Natural language processing; Graph; Artificial intelligence; Multi-document summarization; Information retrieval; Theoretical computer science; Resolution (logic)","score_opus":0.03674720432760291,"score_gpt":0.26818224058725076,"score_spread":0.23143503625964784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2096537984","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09566187,0.0012002084,0.88446075,0.0007201838,0.00015272317,0.00037808993,0.001237089,0.013565312,0.0026238086],"genre_scores_gemma":[0.33391848,0.00071400125,0.65673625,0.00031394447,0.00020953057,0.000258704,0.004708258,0.000761986,0.0023788593],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975496,0.00138296,0.00020806007,0.0003915495,0.00038362487,0.000084137566],"domain_scores_gemma":[0.98761463,0.008677876,0.00081231084,0.0010904766,0.0016735181,0.00013114221],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029094783,0.0009304636,0.0007920102,0.0025900109,0.0006853081,0.0011951637,0.00094424514,0.0009743617,0.0023407862],"category_scores_gemma":[0.015576517,0.0003695303,0.0006544739,0.0024871835,0.00034158345,0.0032706328,0.00088139577,0.0009940052,0.0012187243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008705679,0.00018378731,0.002407124,0.0009026861,0.00027611406,0.00025654474,0.0015119476,0.03653732,0.10012983,0.004938636,0.012992837,0.8389927],"study_design_scores_gemma":[0.00022711507,0.0010150664,0.0048211305,0.00012631674,0.0005859151,0.0004988844,0.0007827423,0.7467353,0.181737,0.030814286,0.032455184,0.00020109791],"about_ca_topic_score_codex":0.0021102983,"about_ca_topic_score_gemma":0.0033489703,"teacher_disagreement_score":0.0029094783,"about_ca_system_score_codex":0.0003901203,"about_ca_system_score_gemma":0.0005166299,"threshold_uncertainty_score":0.015386999},"labels":[],"label_agreement":null},{"id":"W2097120204","doi":"","title":"Towards Robust Abstractive Multi-Document Summarization: A Caseframe Analysis of Centrality and Domain","year":2013,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Automatic summarization; Centrality; Computer science; Redundancy (engineering); Sentence; Natural language processing; Information retrieval; Domain (mathematical analysis); Artificial intelligence; Abstraction; Multi-document summarization","score_opus":0.020833600090225578,"score_gpt":0.26804744871907904,"score_spread":0.24721384862885346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097120204","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040767953,0.0007591442,0.9560688,0.00045788867,0.000024643581,0.000094062605,0.00013764395,0.0004161739,0.0012736627],"genre_scores_gemma":[0.4565983,0.00054351887,0.5402343,0.000089558474,0.00015024506,0.00014029002,0.0007705925,0.00021748456,0.0012556737],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997279,0.0013106664,0.00018278605,0.0005679709,0.00054011546,0.00011949417],"domain_scores_gemma":[0.98500454,0.008888086,0.001572664,0.0018312485,0.0024471579,0.00025624907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035965866,0.00079754536,0.0007488864,0.0049431003,0.000990977,0.0025456802,0.0011394157,0.00109075,0.0012047368],"category_scores_gemma":[0.02187389,0.00039624583,0.0008011211,0.0032397474,0.0012543853,0.004688827,0.0015727977,0.0012358986,0.0004541005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007882637,0.0002837,0.0116812,0.00092258287,0.00028850694,0.00087327766,0.0049879327,0.11355018,0.04971266,0.12081436,0.0065386295,0.6895588],"study_design_scores_gemma":[0.00004944402,0.00021264565,0.005276733,0.00008057498,0.00017902837,0.0003835017,0.001204182,0.84313446,0.023100654,0.11394077,0.0123599395,0.00007798223],"about_ca_topic_score_codex":0.0024256583,"about_ca_topic_score_gemma":0.0024922465,"teacher_disagreement_score":0.0049431003,"about_ca_system_score_codex":0.0009827525,"about_ca_system_score_gemma":0.00085353845,"threshold_uncertainty_score":0.019020796},"labels":[],"label_agreement":null},{"id":"W2100002341","doi":"","title":"Replicated Softmax: an Undirected Topic Model","year":2009,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":442,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Latent Dirichlet allocation; Softmax function; Computer science; Graphical model; Inference; Artificial intelligence; Topic model; Latent variable; Undirected graph; Dirichlet distribution; Theoretical computer science; Data mining; Machine learning; Algorithm; Graph; Deep learning; Mathematics","score_opus":0.03266415990600878,"score_gpt":0.2715287157289674,"score_spread":0.23886455582295862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100002341","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0101081515,0.00019017476,0.9862272,0.00037140818,0.000059358466,0.000055936875,0.00047833667,0.0010164203,0.00149297],"genre_scores_gemma":[0.51639986,0.00069809245,0.46549907,0.0007429844,0.00031084675,0.000798258,0.0027078246,0.0005161922,0.01232681],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99850863,0.00064808363,0.00007517381,0.00038235242,0.00027405965,0.00011161199],"domain_scores_gemma":[0.99778146,0.0013749318,0.00017626025,0.00039731734,0.00019442009,0.0000756003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025367697,0.0010113448,0.0012157746,0.0011165396,0.00043402892,0.0018417579,0.0039252914,0.001801184,0.004651243],"category_scores_gemma":[0.008755062,0.0006968254,0.0014142852,0.0017781013,0.001081588,0.003112356,0.002037663,0.0027897723,0.002248765],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007663705,0.00031134905,0.002839805,0.0003650409,0.00033393115,0.0003258665,0.00036707512,0.5802077,0.0139439255,0.10083622,0.013609709,0.286093],"study_design_scores_gemma":[0.000016015765,0.000022631884,0.00023375821,0.00001131884,0.000023032897,0.000040812483,0.000009474785,0.9637141,0.0015561633,0.032825094,0.0015324902,0.000015133978],"about_ca_topic_score_codex":0.002743036,"about_ca_topic_score_gemma":0.0038698155,"teacher_disagreement_score":0.004651243,"about_ca_system_score_codex":0.0009531957,"about_ca_system_score_gemma":0.0009476447,"threshold_uncertainty_score":0.015559971},"labels":[],"label_agreement":null},{"id":"W2100958968","doi":"10.1109/ccece.2007.310","title":"Speeding Up QA: An Index Structure for Question Queries","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Search engine indexing; Inverted index; Computer science; Bottleneck; Index (typography); Search engine; Scalability; Information retrieval; Question answering; Position (finance); Space (punctuation); Simple (philosophy); Disadvantage; Data mining; Artificial intelligence; Database; World Wide Web","score_opus":0.027307777981333346,"score_gpt":0.3046463902666466,"score_spread":0.27733861228531326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100958968","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011568887,0.001083879,0.9686359,0.0008301481,0.0002640332,0.00072086044,0.0018878522,0.010209594,0.0047988505],"genre_scores_gemma":[0.067090645,0.0008263437,0.91983646,0.0003910115,0.0004334187,0.00074226136,0.005009056,0.0008963296,0.004774504],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99844116,0.00033923704,0.00027778838,0.00024034322,0.0005763533,0.00012513061],"domain_scores_gemma":[0.992961,0.002308911,0.0005338418,0.0019935821,0.0018928483,0.0003097537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023234244,0.0008768774,0.001405318,0.005378399,0.0017502968,0.0020807043,0.0021187882,0.0013780699,0.007227287],"category_scores_gemma":[0.0120379245,0.00078181404,0.0010422519,0.005530248,0.001063865,0.0080292635,0.0032378542,0.0017326022,0.0063717584],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007560099,0.0004286022,0.0020327785,0.00072003016,0.00007160772,0.00023753043,0.0013396667,0.0080329655,0.04381353,0.06499261,0.06373781,0.8138368],"study_design_scores_gemma":[0.0005491916,0.0013331276,0.0030618845,0.0003633719,0.00033363342,0.0016402989,0.00081794005,0.4191064,0.06736127,0.18710609,0.3179617,0.00036517018],"about_ca_topic_score_codex":0.003946904,"about_ca_topic_score_gemma":0.004177449,"teacher_disagreement_score":0.007227287,"about_ca_system_score_codex":0.0010725469,"about_ca_system_score_gemma":0.002101086,"threshold_uncertainty_score":0.02417773},"labels":[],"label_agreement":null},{"id":"W2101916644","doi":"","title":"The Alyssa System at TAC QA 2008","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Component (thermodynamics); Track (disk drive); Entertainment; Information retrieval; Physics; Operating system; Law; Political science","score_opus":0.011242958268532133,"score_gpt":0.22230698129982082,"score_spread":0.21106402303128868,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2101916644","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17321913,0.0037180344,0.2183142,0.005571667,0.0022061171,0.0026739347,0.07373334,0.46198413,0.058579464],"genre_scores_gemma":[0.4798637,0.00071831903,0.2573322,0.0024665105,0.0008170796,0.0012366448,0.19940756,0.009081378,0.049076613],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99639255,0.0013236782,0.00025483404,0.000773741,0.0009931909,0.00026203968],"domain_scores_gemma":[0.9940514,0.0013733857,0.0001827125,0.0010917574,0.0028019005,0.0004987559],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006201586,0.001064742,0.0012509038,0.0024051494,0.0014965978,0.0031415722,0.0022079109,0.0020418803,0.030036261],"category_scores_gemma":[0.009195809,0.00071595976,0.00076328724,0.0017671931,0.0007589076,0.004457084,0.0022664594,0.002050207,0.018449083],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029050023,0.001089789,0.0051338235,0.0013530744,0.00026060257,0.00050696643,0.0018329711,0.012067191,0.033844072,0.008364448,0.6696031,0.26303902],"study_design_scores_gemma":[0.0024128417,0.0011875913,0.009891707,0.00022485707,0.0003071325,0.0006467786,0.00096387696,0.33925918,0.05444393,0.013373182,0.57694745,0.00034150007],"about_ca_topic_score_codex":0.024243414,"about_ca_topic_score_gemma":0.015804574,"teacher_disagreement_score":0.030036261,"about_ca_system_score_codex":0.0019407036,"about_ca_system_score_gemma":0.002048513,"threshold_uncertainty_score":0.10048133},"labels":[],"label_agreement":null},{"id":"W2104483432","doi":"","title":"Probabilistic Document Modeling for Syntax Removal in Text Summarization","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Automatic summarization; Computer science; Probabilistic logic; Syntax; Artificial intelligence; Natural language processing; Metric (unit); Statistical model; Term (time); Generative model; Hidden Markov model; Generative grammar; A priori and a posteriori; Information retrieval","score_opus":0.05544849043075354,"score_gpt":0.253299286546898,"score_spread":0.19785079611614445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104483432","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066426927,0.00071189034,0.98945695,0.00020860658,0.000044671524,0.000063485226,0.00048449324,0.0020255912,0.0003616091],"genre_scores_gemma":[0.30778858,0.0018407591,0.6765285,0.00028152092,0.00047291606,0.0010005825,0.0066298875,0.00091632147,0.004540856],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977876,0.0010106091,0.00019213789,0.00040735686,0.00050017284,0.00010221001],"domain_scores_gemma":[0.99366575,0.004435835,0.0005422155,0.00049188535,0.0007805848,0.00008371797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029965818,0.0010988007,0.0015386487,0.002895547,0.0006965771,0.0016358894,0.0019158972,0.0013179535,0.0016614808],"category_scores_gemma":[0.010376438,0.0007203844,0.0016349715,0.0031337987,0.00053206243,0.0029412152,0.00090505404,0.0016069387,0.002080936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002893811,0.00018167052,0.0020581381,0.0005248237,0.00029294044,0.00018569053,0.00049963564,0.5320758,0.011179299,0.023975879,0.01269864,0.4160382],"study_design_scores_gemma":[0.000015131361,0.00003072427,0.000267666,0.000013764627,0.000030466133,0.000046206194,0.000020141704,0.98357725,0.0013972385,0.012631813,0.0019493129,0.000020259818],"about_ca_topic_score_codex":0.0058809714,"about_ca_topic_score_gemma":0.007867199,"teacher_disagreement_score":0.0058809714,"about_ca_system_score_codex":0.0011464887,"about_ca_system_score_gemma":0.0012063633,"threshold_uncertainty_score":0.015847623},"labels":[],"label_agreement":null},{"id":"W2105382461","doi":"10.1145/1943403.1943452","title":"A reinforcement learning framework for answering complex questions","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Reinforcement learning; Artificial intelligence; Question answering; Support vector machine; Machine learning; Function (biology); Feature engineering; Feature (linguistics); Natural language processing; Deep learning","score_opus":0.10545968433774591,"score_gpt":0.3006082791712535,"score_spread":0.19514859483350758,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105382461","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006401728,0.0002037628,0.99145585,0.00022904424,0.000029777073,0.000071761686,0.000045530887,0.00056488934,0.000997593],"genre_scores_gemma":[0.46081558,0.0003707673,0.53239816,0.000262207,0.00016191414,0.0005754979,0.00027861938,0.00009375456,0.005043591],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992943,0.0003531828,0.000035406047,0.00013430412,0.00013479532,0.000047963393],"domain_scores_gemma":[0.99837637,0.0011169431,0.00011722197,0.0000912581,0.00021081071,0.000087365064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019601795,0.0007219913,0.00085861864,0.00048378584,0.0003707791,0.0006979479,0.0014609025,0.001000431,0.0027524084],"category_scores_gemma":[0.0044720275,0.0003215481,0.0005424561,0.00043557805,0.0007333073,0.0011735522,0.0006533238,0.0013048496,0.0005934329],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020655953,0.0003155157,0.0010742983,0.00017725184,0.00008950124,0.00017635466,0.00026942397,0.7206183,0.0049158065,0.04254221,0.0043145483,0.2253002],"study_design_scores_gemma":[0.000019577044,0.00004068259,0.00007580225,0.0000045053916,0.000005753001,0.00001638057,0.0000071948043,0.98534596,0.00044349334,0.013098017,0.00093660306,0.0000060882253],"about_ca_topic_score_codex":0.003727144,"about_ca_topic_score_gemma":0.004434546,"teacher_disagreement_score":0.003727144,"about_ca_system_score_codex":0.0009211896,"about_ca_system_score_gemma":0.0009875642,"threshold_uncertainty_score":0.010366499},"labels":[],"label_agreement":null},{"id":"W2106344701","doi":"10.1145/1871437.1871556","title":"Automatically suggesting topics for augmenting text documents","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Complement (music); Ranking (information retrieval); Information retrieval; Context (archaeology); Key (lock); Natural language processing; Artificial intelligence","score_opus":0.016657930866883428,"score_gpt":0.277368514942678,"score_spread":0.26071058407579456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106344701","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018687049,0.0012216227,0.962472,0.00041555671,0.0002419377,0.00033256918,0.00085274834,0.013356815,0.0024195882],"genre_scores_gemma":[0.07130361,0.0005502005,0.92153525,0.000095605596,0.00028835417,0.0004481453,0.001790582,0.00092426315,0.0030640494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99774855,0.0006972528,0.00016918957,0.00071430893,0.00055971625,0.00011104181],"domain_scores_gemma":[0.99244833,0.004397994,0.0004874967,0.0009528733,0.0015501407,0.0001631566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020557968,0.0023596012,0.001204949,0.0067271586,0.0014733046,0.0022887243,0.0017463131,0.0017260883,0.004941119],"category_scores_gemma":[0.012663739,0.0009755972,0.0012284116,0.00433718,0.0007135879,0.0042696693,0.0016150214,0.0015473547,0.0043384153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042452035,0.0002914025,0.004297736,0.0011780183,0.00015855295,0.0004888458,0.0022733733,0.014239919,0.052356973,0.0086977845,0.026460074,0.8891328],"study_design_scores_gemma":[0.0002041031,0.00044211402,0.0047026537,0.00039691967,0.0005572734,0.0016036241,0.0011464132,0.6403437,0.1330266,0.042827636,0.174381,0.00036784765],"about_ca_topic_score_codex":0.001892287,"about_ca_topic_score_gemma":0.0041973423,"teacher_disagreement_score":0.0067271586,"about_ca_system_score_codex":0.00053210655,"about_ca_system_score_gemma":0.0015155257,"threshold_uncertainty_score":0.01652962},"labels":[],"label_agreement":null},{"id":"W2107134232","doi":"10.1109/icdmw.2006.71","title":"Enhancing Text Retrieval Performance using Conceptual Ontological Graph","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language processing; Phrase; Sentence; Text graph; Semantics (computer science); Search engine indexing; Representation (politics); Graph; Information retrieval; Artificial intelligence; Precision and recall; Term (time); Conceptual graph; Knowledge representation and reasoning; Theoretical computer science; Text mining","score_opus":0.03618108570638048,"score_gpt":0.2501687320287473,"score_spread":0.2139876463223668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107134232","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.111392155,0.0005384025,0.8784061,0.00051234424,0.00004904635,0.00021713776,0.0004586829,0.0053596348,0.0030664613],"genre_scores_gemma":[0.4648661,0.0007488813,0.5298671,0.0001653976,0.000053876993,0.00022985456,0.0020319677,0.00025451134,0.0017822887],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993455,0.00026916835,0.00004920524,0.00008444691,0.00020503631,0.00004667989],"domain_scores_gemma":[0.99835616,0.00097632897,0.000106552296,0.00026250174,0.00026981148,0.000028681172],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00104571,0.000594941,0.0008190797,0.0025207803,0.00031638923,0.0011518357,0.0010188236,0.00075673027,0.0012787313],"category_scores_gemma":[0.0068564843,0.00014795834,0.0006364523,0.0032443025,0.00029436243,0.0033919008,0.00092043425,0.00043172162,0.00067603413],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006598678,0.00034748123,0.0019011676,0.00066089426,0.000121499055,0.00019568275,0.0004931207,0.081087895,0.07845563,0.028260615,0.00799158,0.79982454],"study_design_scores_gemma":[0.00013018136,0.0003029851,0.001715279,0.000034526784,0.00016565285,0.000355884,0.00027256168,0.9202631,0.03910204,0.026610376,0.0109733995,0.00007400263],"about_ca_topic_score_codex":0.003564663,"about_ca_topic_score_gemma":0.0033191454,"teacher_disagreement_score":0.003564663,"about_ca_system_score_codex":0.0006881964,"about_ca_system_score_gemma":0.0007229769,"threshold_uncertainty_score":0.0070878267},"labels":[],"label_agreement":null},{"id":"W2109209896","doi":"","title":"Measuring Lexical Cohesion: Beyond Word Repetition","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Cohesion (chemistry); Computer science; Vocabulary; Natural language processing; Antecedent (behavioral psychology); Noun phrase; Pronoun; Artificial intelligence; Linguistics; Phrase; Expression (computer science); Personal pronoun; Repetition (rhetorical device); Noun; Metric (unit); Psychology; Social psychology","score_opus":0.04119341609654793,"score_gpt":0.2294941727572702,"score_spread":0.18830075666072227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109209896","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68879306,0.0045720465,0.28394687,0.0005218628,0.0002162297,0.00060205173,0.0034544303,0.0028928057,0.015000674],"genre_scores_gemma":[0.9051924,0.00065306114,0.08836259,0.00007967021,0.00021253231,0.00034253788,0.0031051715,0.0004335779,0.0016184319],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9926513,0.0024235435,0.0010277996,0.001739075,0.0017748728,0.0003833981],"domain_scores_gemma":[0.96269983,0.023125354,0.005768736,0.0036624246,0.003954455,0.0007890862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038988763,0.0010356496,0.0020857826,0.015145047,0.001742219,0.003945345,0.0013732042,0.0015279577,0.0032461486],"category_scores_gemma":[0.04706767,0.0006250696,0.0010338717,0.015001291,0.0013909672,0.009507215,0.0046049207,0.0011961482,0.0017150237],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001166137,0.00044121817,0.21233217,0.0029270835,0.0011248316,0.0009192891,0.01763874,0.014787636,0.03916064,0.017133908,0.006260303,0.686108],"study_design_scores_gemma":[0.00026831988,0.0029660058,0.47276425,0.000949068,0.0017183704,0.003874951,0.021820558,0.24109264,0.05113967,0.14763962,0.054797873,0.00096872303],"about_ca_topic_score_codex":0.003661156,"about_ca_topic_score_gemma":0.0036254197,"teacher_disagreement_score":0.015145047,"about_ca_system_score_codex":0.0008259475,"about_ca_system_score_gemma":0.0011676673,"threshold_uncertainty_score":0.020619452},"labels":[],"label_agreement":null},{"id":"W2109830295","doi":"10.1162/coli.2006.32.3.379","title":"Similarity of Semantic Relations","year":2006,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":179,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"University of Waterloo; Johns Hopkins University","keywords":"Analogy; Word (group theory); Similarity (geometry); Vector space model; Semantic similarity; Relation (database); Vector space; Latent semantic analysis","score_opus":0.01936146345015635,"score_gpt":0.25841519200576835,"score_spread":0.239053728555612,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109830295","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33396602,0.010426541,0.48856002,0.0029460213,0.0008200454,0.0011159485,0.011610308,0.0016962795,0.14885877],"genre_scores_gemma":[0.8810795,0.0021100338,0.10011648,0.00042121042,0.0003368448,0.00064722233,0.009251874,0.00020614309,0.005830756],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9922525,0.0021548334,0.0009495632,0.0016153692,0.0027121536,0.00031548904],"domain_scores_gemma":[0.99073505,0.0047758413,0.0009562263,0.0017220479,0.0015578436,0.00025305766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002787363,0.00066949276,0.0007959002,0.010898451,0.001171668,0.0049129366,0.0011289772,0.0012224552,0.010314451],"category_scores_gemma":[0.029747467,0.000333848,0.0014021558,0.007791879,0.002504809,0.009063483,0.003523278,0.0011885188,0.0021486098],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004952037,0.0002559614,0.034244165,0.0014639018,0.00079078303,0.0005878806,0.007331467,0.0061908974,0.012394395,0.53222245,0.012750557,0.3912723],"study_design_scores_gemma":[0.00005771829,0.00020781964,0.029556204,0.00029817404,0.00025562083,0.001243078,0.0041926648,0.024877891,0.003406513,0.8615936,0.07420375,0.000106984546],"about_ca_topic_score_codex":0.0014244764,"about_ca_topic_score_gemma":0.0010457778,"teacher_disagreement_score":0.010898451,"about_ca_system_score_codex":0.0014667349,"about_ca_system_score_gemma":0.0009266726,"threshold_uncertainty_score":0.034505308},"labels":[],"label_agreement":null},{"id":"W2113673066","doi":"10.1109/hicss.2008.129","title":"Document Retrieval Using Proximity-Based Phrase Searching","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Phrase; Computer science; Section (typography); Information retrieval; Vector space model; Phrase search; Space (punctuation); Matching (statistics); Natural language processing; Artificial intelligence; Search engine; Web search query; Mathematics","score_opus":0.06451732651132608,"score_gpt":0.29511614081772775,"score_spread":0.23059881430640167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2113673066","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038939536,0.0020490314,0.9485933,0.0001542773,0.00009730574,0.00049923983,0.00053579087,0.004427772,0.0047037746],"genre_scores_gemma":[0.2557648,0.0012044184,0.7365171,0.00011858468,0.0002396404,0.00037323273,0.0012947124,0.00019750878,0.0042900634],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978333,0.00053410186,0.00022771095,0.00042642557,0.000859366,0.00011903338],"domain_scores_gemma":[0.99759394,0.0012235014,0.00025562226,0.0004455505,0.00041839233,0.00006300673],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017874032,0.0006540954,0.0016501421,0.004941037,0.000875283,0.0017886049,0.0011280983,0.001171269,0.0053603435],"category_scores_gemma":[0.007351538,0.00037211,0.00083845865,0.006159962,0.00065544737,0.0034784093,0.0017289288,0.0006019794,0.004849879],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009128754,0.0002646683,0.0021939664,0.00077245355,0.00020226225,0.00037227434,0.00053237163,0.018101651,0.09091999,0.018201087,0.009811529,0.8577149],"study_design_scores_gemma":[0.00052947213,0.001852541,0.0061594704,0.0001292388,0.00032943513,0.004020183,0.00055622065,0.7371859,0.15054803,0.05664087,0.0416837,0.0003648814],"about_ca_topic_score_codex":0.0023100742,"about_ca_topic_score_gemma":0.0017850618,"teacher_disagreement_score":0.0053603435,"about_ca_system_score_codex":0.00045785966,"about_ca_system_score_gemma":0.00077101984,"threshold_uncertainty_score":0.017932117},"labels":[],"label_agreement":null},{"id":"W2114562718","doi":"10.1613/jair.2784","title":"Complex Question Answering: Unsupervised Learning Approaches and Experiments","year":2009,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; University of Lethbridge","keywords":"Automatic summarization; Set (abstract data type); Tree kernel; Weighting; Relevance (law); Unsupervised learning; ENCODE; Kernel (algebra); Feature (linguistics)","score_opus":0.4485195580979168,"score_gpt":0.4474377941188779,"score_spread":0.0010817639790389189,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114562718","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7591165,0.0021453598,0.21022129,0.00090575154,0.00029287944,0.0029320642,0.0055439156,0.005009597,0.01383272],"genre_scores_gemma":[0.7684226,0.0006208124,0.21160853,0.00056642847,0.00016206446,0.0020772046,0.012088141,0.00045138106,0.0040029394],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99306715,0.0037110113,0.00050669,0.0012603789,0.0011357432,0.0003190203],"domain_scores_gemma":[0.9679272,0.024277663,0.00095005025,0.0040332507,0.00219508,0.0006166896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007609872,0.0013053673,0.0013219637,0.001293766,0.0010236192,0.0011280854,0.002311356,0.0022194923,0.0037268784],"category_scores_gemma":[0.024734033,0.0004371383,0.0009391369,0.0019255403,0.0014577362,0.0024910388,0.0017798879,0.002207145,0.001511545],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044598193,0.015839985,0.022639483,0.0028209556,0.0010336356,0.00082663354,0.0024632295,0.401255,0.018496336,0.009668214,0.032425214,0.4880715],"study_design_scores_gemma":[0.00061585795,0.001282715,0.009758797,0.00009425888,0.00011657504,0.00045330427,0.00061781716,0.94869447,0.015228538,0.014619494,0.008398741,0.00011933897],"about_ca_topic_score_codex":0.005965608,"about_ca_topic_score_gemma":0.0051007816,"teacher_disagreement_score":0.007609872,"about_ca_system_score_codex":0.0012040398,"about_ca_system_score_gemma":0.0010850427,"threshold_uncertainty_score":0.040245354},"labels":[],"label_agreement":null},{"id":"W2115008472","doi":"10.1162/tacl_a_00233","title":"Distributional Semantics Beyond Words: Supervised Learning of Analogy and Paraphrase","year":2013,"lang":"en","type":"preprint","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Paraphrase; Natural language processing; Similarity (geometry); Artificial intelligence; Computer science; Analogy; Distributional semantics; Word (group theory); Tuple; Pairwise comparison; Noun; Function (biology); Noun phrase; Semantic similarity; Linguistics; Mathematics","score_opus":0.016507278521685844,"score_gpt":0.25692435607696473,"score_spread":0.24041707755527889,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115008472","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2442998,0.0023685512,0.7426266,0.0011759911,0.00012292324,0.00023156549,0.00076297234,0.0016345155,0.0067769815],"genre_scores_gemma":[0.91325074,0.00033675233,0.08240594,0.00027273415,0.00018124229,0.00020358738,0.0019357933,0.00010310715,0.00131002],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959468,0.002221786,0.0002515734,0.0009380958,0.0005056971,0.00013610342],"domain_scores_gemma":[0.98967,0.006453741,0.00094027724,0.0017052173,0.00086976273,0.00036101168],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032841144,0.00091714994,0.0014385148,0.0031885204,0.0008811745,0.0017626355,0.0021556395,0.0017088449,0.0016635197],"category_scores_gemma":[0.018825782,0.00039118985,0.0011405298,0.0028639976,0.0014945164,0.0064441417,0.0027774058,0.0025023846,0.00067861524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008186728,0.0014080633,0.029062033,0.0006459027,0.00055563275,0.00039964655,0.0018264543,0.094848655,0.0059104227,0.052333187,0.012429037,0.79976225],"study_design_scores_gemma":[0.00007330934,0.00021573999,0.0048825517,0.000052793577,0.000055269396,0.00019352531,0.0003051455,0.7836914,0.0016947786,0.20654543,0.0022494507,0.000040634106],"about_ca_topic_score_codex":0.0010844897,"about_ca_topic_score_gemma":0.0015845629,"teacher_disagreement_score":0.0032841144,"about_ca_system_score_codex":0.0007890144,"about_ca_system_score_gemma":0.0007073039,"threshold_uncertainty_score":0.017368257},"labels":[],"label_agreement":null},{"id":"W2115284802","doi":"10.3115/1667583.1667617","title":"Query-focused summaries or query-biased summaries?","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo","keywords":"Computer science; Automatic summarization; Query expansion; Information retrieval; Web query classification; Focus (optics); Query optimization; Web search query; Task (project management); Context (archaeology); Query language; RDF query language; Term (time); Sargable; Multi-document summarization; Search engine","score_opus":0.03382357611830247,"score_gpt":0.25941036811570356,"score_spread":0.2255867919974011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115284802","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19715618,0.06600941,0.63716,0.048210528,0.0025065243,0.00075420603,0.0062579447,0.011439341,0.030505804],"genre_scores_gemma":[0.808773,0.007632464,0.16430615,0.0051695798,0.0030199073,0.000335053,0.0029144492,0.00090714375,0.0069422135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9947277,0.0024531314,0.0004155249,0.0010071219,0.0011480136,0.00024854613],"domain_scores_gemma":[0.9720958,0.015273227,0.0034283176,0.0045074825,0.0041619474,0.000533269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008836453,0.00083535886,0.0012741679,0.0020270352,0.000475548,0.002838358,0.0011388025,0.0020436242,0.0054526716],"category_scores_gemma":[0.0659209,0.00040127133,0.0002996292,0.0024831423,0.0010295964,0.008295595,0.0010991603,0.0008958602,0.0024566432],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018948147,0.0002236743,0.013273552,0.0025331085,0.00037368483,0.00046402914,0.0033614929,0.006397617,0.018900534,0.08888749,0.070341386,0.7933488],"study_design_scores_gemma":[0.0007635377,0.0020437741,0.032360524,0.0013628483,0.0008425028,0.0042631496,0.0051261187,0.13541003,0.06203735,0.49131212,0.2639893,0.0004887377],"about_ca_topic_score_codex":0.0012861473,"about_ca_topic_score_gemma":0.0014340974,"teacher_disagreement_score":0.008836453,"about_ca_system_score_codex":0.00077821646,"about_ca_system_score_gemma":0.0008162537,"threshold_uncertainty_score":0.046732187},"labels":[],"label_agreement":null},{"id":"W2115791615","doi":"","title":"Deep Learning for NLP (without Magic)","year":2012,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":132,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Feature engineering; Deep learning; Artificial neural network; Paraphrase; Machine learning; Sentiment analysis; Natural language processing; Feature (linguistics); Focus (optics); MAGIC (telescope); Language model","score_opus":0.022720816901939583,"score_gpt":0.2786095360606215,"score_spread":0.25588871915868194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115791615","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011612065,0.032425556,0.9131408,0.009356223,0.0024057762,0.00009755156,0.0008143042,0.0036204022,0.036978193],"genre_scores_gemma":[0.046705414,0.06686021,0.77927786,0.006987286,0.0039438056,0.00083610957,0.002979167,0.0021894865,0.090220705],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9993895,0.00016211112,0.000047655536,0.00011844681,0.00023200714,0.000050234914],"domain_scores_gemma":[0.99933475,0.0003868128,0.00003873346,0.00008929199,0.00011687104,0.000033512675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009801683,0.0011946372,0.0005381608,0.00091861363,0.00048439155,0.0027825877,0.001342165,0.0017726759,0.021926563],"category_scores_gemma":[0.0040751896,0.0006503332,0.0008742206,0.0012362924,0.001296812,0.005107338,0.0019384716,0.0042274785,0.010011784],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039942806,0.000048427737,0.00014878233,0.0010064798,0.00004481854,0.00013761809,0.00017632873,0.014207391,0.0026809787,0.3554519,0.15923527,0.4668221],"study_design_scores_gemma":[0.000016299131,0.000034279143,0.0001663144,0.00047062614,0.000016302813,0.00020670879,0.000045478537,0.052926548,0.0023751718,0.51020676,0.4335012,0.00003428261],"about_ca_topic_score_codex":0.0019265944,"about_ca_topic_score_gemma":0.0021269852,"teacher_disagreement_score":0.021926563,"about_ca_system_score_codex":0.0013482582,"about_ca_system_score_gemma":0.0010135677,"threshold_uncertainty_score":0.07335162},"labels":[],"label_agreement":null},{"id":"W2116535157","doi":"10.1109/ictai.2008.26","title":"Exploiting Syntactic and Shallow Semantic Kernels to Improve Random Walks for Complex Question Answering","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Question answering; Natural language processing; Artificial intelligence; Sentence; Graph; Word (group theory); Information retrieval; Task (project management); Theoretical computer science; Linguistics","score_opus":0.04657806458457651,"score_gpt":0.27570348450718907,"score_spread":0.22912541992261257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116535157","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05454121,0.0003518454,0.9419766,0.00014275072,0.000029120714,0.000059395086,0.000084509265,0.0021369432,0.00067749765],"genre_scores_gemma":[0.5740574,0.0003558817,0.42148307,0.00014542976,0.00008179606,0.000121205,0.0011013331,0.0003976511,0.002256224],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989899,0.0004909032,0.000053682066,0.0002032198,0.00018090711,0.000081351005],"domain_scores_gemma":[0.9949993,0.0034341842,0.00030724774,0.00062588457,0.00046907982,0.00016436803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019894843,0.0009870699,0.0014009952,0.0015539313,0.0005155995,0.0008631628,0.0011752067,0.0012687246,0.0014813669],"category_scores_gemma":[0.009530602,0.0004603085,0.0009802034,0.0013100635,0.0004733859,0.0029366687,0.0011001151,0.0011594782,0.0010390581],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046327707,0.0006174474,0.004096512,0.00025229712,0.00021750116,0.000116070216,0.000340923,0.47645354,0.021838889,0.029469823,0.0035277605,0.46260586],"study_design_scores_gemma":[0.000008257375,0.000035213983,0.00017849347,0.00000258748,0.0000094698535,0.0000088491515,0.0000060828143,0.9902244,0.00080966245,0.008476795,0.0002336684,0.000006462747],"about_ca_topic_score_codex":0.005372953,"about_ca_topic_score_gemma":0.010816893,"teacher_disagreement_score":0.005372953,"about_ca_system_score_codex":0.00063635374,"about_ca_system_score_gemma":0.0007468534,"threshold_uncertainty_score":0.010683358},"labels":[],"label_agreement":null},{"id":"W2117108960","doi":"10.2190/w5ar-dypw-40kx-fl99","title":"Essay Assessment with Latent Semantic Analysis","year":2003,"lang":"en","type":"article","venue":"Journal of Educational Computing Research","topic":"Topic Modeling","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Latent semantic analysis; Automatic summarization; Readability; Computer science; Natural language processing; Semantics (computer science); Semantic similarity; Quality (philosophy); Similarity (geometry); Probabilistic latent semantic analysis; Information retrieval; Artificial intelligence; Semantic analysis (machine learning); Data science","score_opus":0.08015701964176432,"score_gpt":0.4292517564698892,"score_spread":0.3490947368281249,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117108960","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022439152,0.0004583246,0.9718768,0.00032835075,0.00003872293,0.00017524554,0.0003841626,0.0011041687,0.0031950383],"genre_scores_gemma":[0.53849185,0.00050019444,0.4561198,0.00008389711,0.0001501155,0.00054634636,0.0014250389,0.00011846338,0.0025643778],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99126536,0.0054567237,0.00055846025,0.00079636683,0.0017576227,0.00016539286],"domain_scores_gemma":[0.98762983,0.007378144,0.00168877,0.0011867775,0.0019092942,0.00020713187],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005859427,0.0008226928,0.00082833564,0.006551941,0.00060755253,0.0033898028,0.0008295876,0.00072492694,0.0026524186],"category_scores_gemma":[0.030830087,0.00030188178,0.0009837539,0.0039810673,0.0009033349,0.0035941182,0.0024147169,0.0013029061,0.001119911],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027232346,0.00035952838,0.02261651,0.0004731451,0.000442972,0.000121730096,0.0013920314,0.04284103,0.004307573,0.08033375,0.006756205,0.8400832],"study_design_scores_gemma":[0.00006452326,0.00023807026,0.012477243,0.00014001822,0.00011472213,0.00016875479,0.0005832669,0.74769825,0.004725345,0.22438663,0.0092992475,0.00010403515],"about_ca_topic_score_codex":0.0013911049,"about_ca_topic_score_gemma":0.0016215058,"teacher_disagreement_score":0.006551941,"about_ca_system_score_codex":0.0009894582,"about_ca_system_score_gemma":0.0015812946,"threshold_uncertainty_score":0.030987918},"labels":[],"label_agreement":null},{"id":"W2117543632","doi":"10.1145/1367497.1367739","title":"Reasoning about similarity queries in text retrieval tasks","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Similarity (geometry); Information retrieval; Sampling (signal processing); Data mining; Calibration; Document retrieval; Query expansion; Artificial intelligence; Mathematics","score_opus":0.02953565206192772,"score_gpt":0.25511994357873424,"score_spread":0.22558429151680653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117543632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08035501,0.0010092424,0.9139839,0.00085082185,0.000035835354,0.00036460496,0.00061272696,0.0017858883,0.0010019966],"genre_scores_gemma":[0.65897614,0.0007925232,0.33369538,0.00049336883,0.0002952782,0.000538143,0.0038931922,0.0001878249,0.0011282108],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.990267,0.0045145806,0.0011125033,0.0015902853,0.0020694996,0.0004461344],"domain_scores_gemma":[0.9745063,0.020760443,0.0014315388,0.0014046325,0.0014187901,0.00047825588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008626205,0.001388286,0.002670268,0.0036669231,0.0016642641,0.003531452,0.0037083768,0.0034847122,0.0014436528],"category_scores_gemma":[0.0492609,0.0010060966,0.0018376136,0.0040189167,0.0011827237,0.01134416,0.0026763016,0.0024172373,0.0007262339],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028095802,0.001451673,0.016099203,0.0018520121,0.00061160367,0.0015042414,0.004773039,0.40488634,0.025941335,0.0674697,0.015829628,0.4567716],"study_design_scores_gemma":[0.00008446869,0.00010950582,0.0011712204,0.000027024887,0.000052459203,0.00020499414,0.000364717,0.9308471,0.0042368295,0.061665226,0.0011977248,0.000038630817],"about_ca_topic_score_codex":0.00595079,"about_ca_topic_score_gemma":0.0051029283,"teacher_disagreement_score":0.008626205,"about_ca_system_score_codex":0.0017422464,"about_ca_system_score_gemma":0.0012343897,"threshold_uncertainty_score":0.045620263},"labels":[],"label_agreement":null},{"id":"W2120520504","doi":"10.5815/ijmecs.2015.01.01","title":"Semantic Question Generation Using Artificial Immunity","year":2015,"lang":"en","type":"article","venue":"International Journal of Modern Education and Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Sentence; Preprocessor; Artificial intelligence; Natural language processing; Classifier (UML); Set (abstract data type); Test set; Matching (statistics); Semantic role labeling","score_opus":0.08427210239310003,"score_gpt":0.3489781839158917,"score_spread":0.2647060815227917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120520504","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0923354,0.0004324156,0.8953041,0.0012150978,0.00016488452,0.0003457346,0.00027556007,0.0030436239,0.0068832496],"genre_scores_gemma":[0.7824267,0.00020805933,0.21089391,0.00057416456,0.00012411343,0.00028519,0.0010156789,0.00010760685,0.0043644933],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99877256,0.00048633103,0.00007405348,0.00035545163,0.00022751873,0.00008404269],"domain_scores_gemma":[0.99801457,0.0010396614,0.00017837706,0.00021048765,0.0004889459,0.00006792515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012713167,0.00057874975,0.0005107618,0.0011387434,0.00034917152,0.0012263715,0.0010983093,0.0011443833,0.0025805421],"category_scores_gemma":[0.004842102,0.00025719052,0.0012668426,0.00046055586,0.00056757504,0.0024055338,0.0008829529,0.0009955728,0.0009841167],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005929013,0.0010492307,0.011773879,0.0005987597,0.00026748813,0.0007763108,0.002110752,0.16998197,0.07086907,0.06869002,0.014022221,0.65926737],"study_design_scores_gemma":[0.000020781279,0.00013273947,0.0010709408,0.000020167143,0.0000444232,0.00013625997,0.0001027514,0.9589101,0.008251981,0.027037222,0.004251845,0.000020739979],"about_ca_topic_score_codex":0.0013399165,"about_ca_topic_score_gemma":0.0009548222,"teacher_disagreement_score":0.0025805421,"about_ca_system_score_codex":0.0007047612,"about_ca_system_score_gemma":0.0005774719,"threshold_uncertainty_score":0.008632839},"labels":[],"label_agreement":null},{"id":"W2121027694","doi":"","title":"DalTREC 2005 QA System Jellyfish: Mark-and-Match Approach to Question Answering","year":2005,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Dalhousie University","funders":"","keywords":"Question answering; Jellyfish; Rewriting; Computer science; Information retrieval; Natural language processing; Artificial intelligence; Programming language; Ecology; Biology","score_opus":0.016189602661224045,"score_gpt":0.2300128603648973,"score_spread":0.21382325770367325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121027694","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021689402,0.00119819,0.7726822,0.0042334576,0.00066161284,0.0014312373,0.010772552,0.16445261,0.022878788],"genre_scores_gemma":[0.19474474,0.0007008742,0.7071738,0.0017414482,0.00030003977,0.0005695758,0.031245062,0.006699359,0.056825187],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9965432,0.0013406437,0.00025034125,0.0007881793,0.0008258918,0.00025175262],"domain_scores_gemma":[0.9948449,0.0013238449,0.00014379474,0.0015991217,0.0017397547,0.00034860524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064398255,0.0009228697,0.0012430856,0.0021235724,0.0016222964,0.003538239,0.0034334,0.0018245624,0.029164733],"category_scores_gemma":[0.0112958895,0.00084993005,0.00094149134,0.0014908633,0.00089961704,0.0066351607,0.003021294,0.003320379,0.016205303],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010264904,0.0005265313,0.0025965814,0.0010776422,0.00019051359,0.00024335476,0.001599918,0.013389249,0.036116846,0.040794116,0.42400494,0.47843376],"study_design_scores_gemma":[0.0003436091,0.0004734738,0.0027930308,0.000091208385,0.00013373588,0.0005105065,0.0006553246,0.29642683,0.073478214,0.05328377,0.5715645,0.0002458036],"about_ca_topic_score_codex":0.019355115,"about_ca_topic_score_gemma":0.023115527,"teacher_disagreement_score":0.029164733,"about_ca_system_score_codex":0.003000519,"about_ca_system_score_gemma":0.003395849,"threshold_uncertainty_score":0.09756577},"labels":[],"label_agreement":null},{"id":"W2121771196","doi":"","title":"AMBER: A Modified BLEU, Enhanced Ranking Metric","year":2011,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"BLEU; Metric (unit); Computer science; Machine translation; Ranking (information retrieval); Artificial intelligence; Natural language processing; Sentence; Consistency (knowledge bases); Task (project management); Recall; Translation (biology); Speech recognition; Linguistics","score_opus":0.04939886284340537,"score_gpt":0.24351303555473786,"score_spread":0.1941141727113325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121771196","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050654184,0.009815158,0.87355644,0.00082483096,0.0011596743,0.00088353694,0.010703439,0.024111101,0.028291648],"genre_scores_gemma":[0.31450912,0.0016543616,0.64068323,0.00069722603,0.0006220406,0.0015850415,0.022036405,0.0036130964,0.014599434],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98642945,0.0075099003,0.00086835626,0.0011990964,0.003649336,0.000343834],"domain_scores_gemma":[0.9862213,0.0055865822,0.0009782817,0.0028704396,0.0040622964,0.0002811219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007470604,0.0018149354,0.0019329259,0.0071555423,0.0011447186,0.0019482976,0.0023938091,0.0014836517,0.0046264185],"category_scores_gemma":[0.024884954,0.00039343548,0.0007424435,0.0045694634,0.00058899407,0.0032752436,0.0016799091,0.0016349069,0.0044942326],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084350555,0.00038878503,0.0049314033,0.0011310237,0.0006479234,0.00012381672,0.0002806773,0.038187493,0.017901119,0.012524661,0.076093994,0.8469456],"study_design_scores_gemma":[0.00059675786,0.0034719643,0.019624418,0.0003637505,0.000860724,0.0018968345,0.00022257527,0.67294824,0.08248059,0.04863255,0.16795245,0.0009490927],"about_ca_topic_score_codex":0.0029357753,"about_ca_topic_score_gemma":0.0064596147,"teacher_disagreement_score":0.007470604,"about_ca_system_score_codex":0.001164854,"about_ca_system_score_gemma":0.0016118768,"threshold_uncertainty_score":0.03950882},"labels":[],"label_agreement":null},{"id":"W2122963530","doi":"10.7939/r33775w6g","title":"Reliability of Online Scoring of First Mentions in the Edmonton Narrative Normative Instrument","year":2011,"lang":"en","type":"article","venue":"University of Alberta Library","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reliability (semiconductor); Normative; Narrative; Computer science; Reliability engineering; Engineering; Political science; Linguistics; Law","score_opus":0.020423481448785957,"score_gpt":0.18734151343069239,"score_spread":0.16691803198190644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122963530","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96109056,0.0002732492,0.005980701,0.0003246914,0.00021339851,0.00034946957,0.0014418735,0.00028112673,0.030044831],"genre_scores_gemma":[0.9791018,0.00024261785,0.007761454,0.00010801263,0.000055696517,0.0008747684,0.0028793155,0.0001685084,0.008807836],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9835616,0.0070544877,0.0015128624,0.0014172924,0.0058807107,0.0005730927],"domain_scores_gemma":[0.87564856,0.062433187,0.0072223675,0.00868015,0.043818627,0.002197228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028768726,0.00037573828,0.000509448,0.004750763,0.0011222161,0.0023304285,0.0011789909,0.00079081947,0.003292283],"category_scores_gemma":[0.11616766,0.00039867134,0.00062019174,0.0017784939,0.001042173,0.0017754787,0.0026704203,0.0009827005,0.0017757532],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076923746,0.00070424424,0.79869294,0.00029309266,0.00021691431,0.00010878352,0.031848013,0.0013881227,0.0035740656,0.0032394743,0.010453511,0.14871158],"study_design_scores_gemma":[0.000038660655,0.00020908665,0.96308696,0.00028838936,0.00010996687,0.00013988315,0.009147701,0.0033394576,0.0042926883,0.0011027969,0.018157667,0.00008670158],"about_ca_topic_score_codex":0.022955896,"about_ca_topic_score_gemma":0.046467695,"teacher_disagreement_score":0.028768726,"about_ca_system_score_codex":0.0018489963,"about_ca_system_score_gemma":0.0019911039,"threshold_uncertainty_score":0.15214545},"labels":[],"label_agreement":null},{"id":"W2124006395","doi":"10.3758/brm.42.2.393","title":"Exploring lexical co-occurrence space using HiDEx","year":2010,"lang":"en","type":"article","venue":"Behavior Research Methods","topic":"Topic Modeling","field":"Computer Science","cited_by":154,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Representation (politics); Co-occurrence; Set (abstract data type); Space (punctuation); A priori and a posteriori; Semantic space; Lexical decision task; Natural language processing; Artificial intelligence; Basis (linear algebra); Semantic data model; Hyperspace; Mathematics; Cognition","score_opus":0.8360044474353515,"score_gpt":0.6474662355025267,"score_spread":0.18853821193282483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124006395","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30680278,0.0008534086,0.67825013,0.00026839023,0.000058961163,0.00013601239,0.004533009,0.006520036,0.002577222],"genre_scores_gemma":[0.7415824,0.00030403855,0.2483015,0.00003221595,0.00004603372,0.0003547997,0.006355012,0.0005879025,0.0024361596],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989353,0.00048173382,0.00007535297,0.0002767375,0.00013269809,0.000098179305],"domain_scores_gemma":[0.99023,0.008671685,0.00027044167,0.0004721046,0.00020736291,0.0001483653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020279305,0.0007241619,0.00095945754,0.0050021945,0.00081386045,0.0021382102,0.00093934935,0.0006348629,0.008626315],"category_scores_gemma":[0.007404076,0.00038591825,0.0012405773,0.0045088613,0.00041890788,0.0023519683,0.0018518426,0.0008358816,0.0012534198],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028990381,0.0005948304,0.077313654,0.0011913995,0.0008342651,0.00081938994,0.0066940025,0.02418538,0.033079717,0.027517514,0.0082558235,0.8166149],"study_design_scores_gemma":[0.00018751319,0.00035659643,0.039341494,0.000105392915,0.0004017792,0.0007823943,0.0033098834,0.8764722,0.01065702,0.05063141,0.017643142,0.00011132878],"about_ca_topic_score_codex":0.005194644,"about_ca_topic_score_gemma":0.006613588,"teacher_disagreement_score":0.008626315,"about_ca_system_score_codex":0.00037529672,"about_ca_system_score_gemma":0.0008689869,"threshold_uncertainty_score":0.028857887},"labels":[],"label_agreement":null},{"id":"W2125028118","doi":"","title":"AUEB at TAC 2008","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; String (physics); Rank (graph theory); Natural language processing; Artificial intelligence; Support vector machine; Similarity (geometry); Mathematics","score_opus":0.013933346286805605,"score_gpt":0.23478537384350479,"score_spread":0.2208520275566992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125028118","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.118481174,0.006035041,0.12098846,0.039228536,0.0260927,0.0025979644,0.05610132,0.04245382,0.5880209],"genre_scores_gemma":[0.20078717,0.00084991055,0.07657166,0.0035633843,0.0026235087,0.0010772953,0.074170075,0.005527376,0.6348296],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976064,0.000612753,0.000043944787,0.0004437289,0.0007768213,0.00051625905],"domain_scores_gemma":[0.9959578,0.00031760018,0.00007537402,0.00048366716,0.0016493726,0.0015161844],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.005414136,0.00089068786,0.00087256747,0.0012453959,0.0021582309,0.004391448,0.0016277318,0.0012848654,0.09976361],"category_scores_gemma":[0.0056132344,0.00028766503,0.00036958436,0.0007892513,0.00043459883,0.0022320282,0.0022804479,0.0019658315,0.059317414],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007905123,0.00056475034,0.0012541806,0.00007997433,0.000024590508,0.00035486955,0.00039070103,0.0009790944,0.0046158205,0.0058816047,0.86810255,0.116961405],"study_design_scores_gemma":[0.00021487368,0.00027186808,0.0027424288,0.00004876135,0.000017033713,0.00015878264,0.00038936437,0.008729195,0.002900178,0.004523709,0.97997034,0.00003341604],"about_ca_topic_score_codex":0.010486255,"about_ca_topic_score_gemma":0.018437693,"teacher_disagreement_score":0.90023637,"about_ca_system_score_codex":0.0020771332,"about_ca_system_score_gemma":0.00242935,"threshold_uncertainty_score":0.3337425},"labels":[],"label_agreement":null},{"id":"W2125247927","doi":"10.3115/1690219.1690223","title":"Summarizing multiple spoken documents","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Salient; Transcription (linguistics); Speech recognition; Vocabulary; Word (group theory); Spoken language; Word error rate; Artificial intelligence; Natural language processing; Feature (linguistics); Language model; Independence (probability theory); Linguistics","score_opus":0.019453101614236773,"score_gpt":0.25235971961293857,"score_spread":0.2329066179987018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125247927","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032621335,0.0016816553,0.9533317,0.00094012055,0.00052484614,0.00023978639,0.0026577536,0.0033709065,0.004631921],"genre_scores_gemma":[0.42404297,0.0019341025,0.54636115,0.00047374153,0.00090129505,0.0005692203,0.006853488,0.0006204777,0.018243596],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837637,0.00035221688,0.00015042446,0.000547856,0.00048535428,0.000087748245],"domain_scores_gemma":[0.9971988,0.0013262201,0.00028545668,0.00039893534,0.0006978023,0.000092766175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012693905,0.0011341759,0.0010509193,0.0017171851,0.0006105792,0.002241894,0.0017369136,0.001218541,0.0058391606],"category_scores_gemma":[0.006477901,0.00049998,0.0013673544,0.0014357297,0.00047149564,0.0028123318,0.0012678523,0.0012107689,0.0029998275],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007502634,0.00021905989,0.0041389605,0.0013662751,0.00044607575,0.0013623635,0.002899819,0.093708105,0.040671304,0.04306391,0.019307442,0.7920664],"study_design_scores_gemma":[0.00013418142,0.00055002264,0.0041512083,0.00022574433,0.00060175866,0.0013314947,0.0012557927,0.8289921,0.030195694,0.0715449,0.06075857,0.00025845357],"about_ca_topic_score_codex":0.0029116578,"about_ca_topic_score_gemma":0.0031434528,"teacher_disagreement_score":0.0058391606,"about_ca_system_score_codex":0.00070146343,"about_ca_system_score_gemma":0.0009807845,"threshold_uncertainty_score":0.019533992},"labels":[],"label_agreement":null},{"id":"W2126034021","doi":"10.1017/s1351324901002765","title":"Discovery of inference rules for question-answering","year":2001,"lang":"en","type":"article","venue":"Natural Language Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":532,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Question answering; Inference; Computer science; Dependency (UML); Set (abstract data type); Parsing; Artificial intelligence; Natural language processing; Rule of inference; Programming language","score_opus":0.00715054306742324,"score_gpt":0.2517287714730841,"score_spread":0.24457822840566085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126034021","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0067988033,0.0003653209,0.9827832,0.00080081896,0.00005199747,0.0003577637,0.0009942934,0.006214378,0.0016333653],"genre_scores_gemma":[0.05577256,0.0003113596,0.9375631,0.00033869751,0.000121391924,0.00037481636,0.004171626,0.00039183703,0.0009544628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9877671,0.004144249,0.0014878274,0.0032813326,0.002952466,0.00036705678],"domain_scores_gemma":[0.9239963,0.060990427,0.0026269376,0.006629165,0.005140124,0.0006169625],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01285425,0.0016765106,0.001928733,0.007963133,0.0021711406,0.003862006,0.004943471,0.0028762098,0.004924751],"category_scores_gemma":[0.08408172,0.0016174404,0.0034202896,0.003867732,0.002223716,0.009593174,0.0038252957,0.0049922033,0.0038363792],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029533787,0.0006702197,0.014070659,0.0012709078,0.00049688324,0.001019094,0.0034075996,0.026488563,0.011181616,0.11923913,0.028005902,0.7938541],"study_design_scores_gemma":[0.00008116165,0.00005574021,0.0024393387,0.00024227216,0.00022057795,0.0006252131,0.0005403052,0.5626093,0.014832978,0.3893543,0.02891343,0.000085374355],"about_ca_topic_score_codex":0.0038652981,"about_ca_topic_score_gemma":0.0053856303,"teacher_disagreement_score":0.01285425,"about_ca_system_score_codex":0.0016318441,"about_ca_system_score_gemma":0.0026073796,"threshold_uncertainty_score":0.06798059},"labels":[],"label_agreement":null},{"id":"W2127099514","doi":"","title":"Large-Scale Learning of Embeddings with Reconstruction Sampling","year":2011,"lang":"en","type":"article","venue":"International Conference on Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Machine learning; Vocabulary; Estimator; Context (archaeology); Sampling (signal processing); Artificial neural network; Scale (ratio); Encoder; Computer vision; Mathematics","score_opus":0.05268136800075558,"score_gpt":0.28672302850387354,"score_spread":0.23404166050311798,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2127099514","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007837283,0.00010315737,0.99068403,0.00010840119,0.000020955293,0.000029865903,0.000049243317,0.0008926305,0.000274448],"genre_scores_gemma":[0.2696874,0.00029057183,0.7245209,0.00029024782,0.00016257471,0.00037236657,0.0014532813,0.000359377,0.0028632968],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998131,0.00074599683,0.000107563115,0.00045187958,0.00044109664,0.00012249523],"domain_scores_gemma":[0.99326295,0.004112736,0.00036939367,0.001421269,0.0007003625,0.00013330524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002770897,0.0013879938,0.0014370956,0.00094243506,0.0005260048,0.0011722891,0.002139457,0.0016096276,0.0019184566],"category_scores_gemma":[0.019152125,0.000977012,0.0009217142,0.001114142,0.0012073697,0.0045730155,0.0026948573,0.002515605,0.0012432229],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044805827,0.00024837078,0.0034483874,0.0003172345,0.00015949909,0.00027677655,0.0003242351,0.5476606,0.015499743,0.049280237,0.008650067,0.3736868],"study_design_scores_gemma":[0.000022970038,0.000033452587,0.00014780356,0.000005981639,0.000007249652,0.0000373682,0.000018084478,0.98123217,0.0026513597,0.015017941,0.0008175885,0.000007891834],"about_ca_topic_score_codex":0.002501535,"about_ca_topic_score_gemma":0.003958491,"teacher_disagreement_score":0.002770897,"about_ca_system_score_codex":0.00076125,"about_ca_system_score_gemma":0.001124028,"threshold_uncertainty_score":0.0146541},"labels":[],"label_agreement":null},{"id":"W2128424286","doi":"10.1017/s1351324911000167","title":"Query-focused multi-document summarization: automatic data annotations and supervised learning approaches","year":2011,"lang":"en","type":"article","venue":"Natural Language Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Automatic summarization; Artificial intelligence; Conditional random field; Semi-supervised learning; Supervised learning; Support vector machine; Annotation; Machine learning; Information retrieval; Natural language processing; Artificial neural network","score_opus":0.05869544911709373,"score_gpt":0.23725411456844478,"score_spread":0.17855866545135105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128424286","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045954797,0.0012153274,0.9453755,0.0003531661,0.00006171583,0.00020563371,0.00031595657,0.0057233577,0.00079456466],"genre_scores_gemma":[0.28471115,0.00042646716,0.7108411,0.00015304884,0.00018756583,0.00026469698,0.001936125,0.00030608964,0.0011738337],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946114,0.0030173871,0.00037861546,0.0009843139,0.0008850521,0.00012321462],"domain_scores_gemma":[0.9788084,0.011889986,0.0022405484,0.0027872499,0.003968779,0.00030505154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053687473,0.001290555,0.0013351663,0.0033314594,0.0008739584,0.0012193682,0.0016706086,0.0011543587,0.0007106849],"category_scores_gemma":[0.015067327,0.0004132582,0.0009309745,0.002231134,0.00059046806,0.002876069,0.0010589629,0.0013578049,0.0007266281],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005805876,0.00073128566,0.0038840824,0.0006151026,0.00032027016,0.00012340664,0.00062352355,0.07474454,0.02671475,0.0021542676,0.004666821,0.8848415],"study_design_scores_gemma":[0.00006551621,0.00034503042,0.0030827601,0.000056201527,0.00016547575,0.00011343441,0.00020584618,0.9503317,0.035996318,0.0060813823,0.0034919474,0.00006445294],"about_ca_topic_score_codex":0.0021132242,"about_ca_topic_score_gemma":0.0043288297,"teacher_disagreement_score":0.0053687473,"about_ca_system_score_codex":0.0007236545,"about_ca_system_score_gemma":0.00090369175,"threshold_uncertainty_score":0.02839303},"labels":[],"label_agreement":null},{"id":"W2128739123","doi":"10.1609/aaai.v28i1.8933","title":"Learning Concept Embeddings for Query Expansion by Quantum Entropy Minimization","year":2014,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Query expansion; Computer science; Information retrieval; Embedding; Web search query; Representation (politics); Sargable; Web query classification; Theoretical computer science; Search engine; Artificial intelligence","score_opus":0.04401282363568374,"score_gpt":0.2829489593546883,"score_spread":0.23893613571900454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128739123","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023220662,0.00037335142,0.97490644,0.00027418692,0.000023098464,0.000050234146,0.00005968196,0.000253698,0.0008386119],"genre_scores_gemma":[0.71061355,0.0007568438,0.28310037,0.00045402942,0.00019313768,0.000372152,0.0005974804,0.00024318915,0.0036693113],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998965,0.00049417483,0.00005531559,0.00018030668,0.00022900991,0.00007610319],"domain_scores_gemma":[0.99731076,0.0020710311,0.00015112718,0.00019678545,0.00020173604,0.0000686176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019366579,0.00070725934,0.001032321,0.00080944767,0.00033527726,0.00077515235,0.0011680952,0.0010255787,0.0017227664],"category_scores_gemma":[0.007358598,0.0004547035,0.0007963725,0.0009748088,0.0012497754,0.0032741854,0.0017527062,0.001738798,0.00042402645],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037335977,0.00036253713,0.0014902578,0.0003451298,0.00012274383,0.00015016017,0.00039920083,0.5692148,0.012656911,0.15419716,0.0067432574,0.25394434],"study_design_scores_gemma":[0.00001045364,0.0000353775,0.000086108,0.00000426065,0.0000055008063,0.000016933775,0.000011125804,0.9693839,0.000535237,0.029616281,0.00028780734,0.000006990567],"about_ca_topic_score_codex":0.0013768977,"about_ca_topic_score_gemma":0.0014100006,"teacher_disagreement_score":0.0019366579,"about_ca_system_score_codex":0.0008768055,"about_ca_system_score_gemma":0.00070156774,"threshold_uncertainty_score":0.010242164},"labels":[],"label_agreement":null},{"id":"W2129198817","doi":"10.1109/icdmw.2009.49","title":"Parameterized Contrast in Second Order Soft Co-occurrences: A Novel Text Representation Technique in Text Mining and Knowledge Extraction","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Contrast (vision); Representation (politics); Natural language processing; Artificial intelligence; Knowledge extraction; Meaning (existential); Parameterized complexity; Information extraction; Information retrieval; Co-occurrence; Relationship extraction; Word (group theory); Linguistics; Algorithm","score_opus":0.04513540965388924,"score_gpt":0.3360023972293884,"score_spread":0.29086698757549917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129198817","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018186443,0.00053192535,0.97885394,0.0002365784,0.00006170516,0.00011140101,0.0006093567,0.0006971603,0.0007115539],"genre_scores_gemma":[0.2629285,0.00045773695,0.7318332,0.0001812234,0.00024928426,0.0006601068,0.0021937203,0.00025388974,0.0012423779],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961682,0.001147782,0.00050291343,0.0009950007,0.001034349,0.00015180073],"domain_scores_gemma":[0.98543173,0.010624914,0.0013721209,0.0012654115,0.0010919957,0.00021380615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031439145,0.00076028146,0.0011071169,0.0072923494,0.0007280442,0.0022371386,0.0012490347,0.0011989087,0.0021463824],"category_scores_gemma":[0.018589256,0.0004399002,0.0012620054,0.009014906,0.001043212,0.0052103763,0.0015478624,0.0018474548,0.0009927276],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007549113,0.00025600195,0.010441717,0.0009164806,0.0003432002,0.0005978415,0.001365414,0.023567582,0.039328646,0.080770254,0.0059081847,0.8357498],"study_design_scores_gemma":[0.000093339964,0.00036829463,0.009499821,0.00018743676,0.00018728382,0.0017319493,0.00046801803,0.7744402,0.026393961,0.1602969,0.02613793,0.00019489725],"about_ca_topic_score_codex":0.0009676006,"about_ca_topic_score_gemma":0.0013742527,"teacher_disagreement_score":0.0072923494,"about_ca_system_score_codex":0.00077557244,"about_ca_system_score_gemma":0.0013406645,"threshold_uncertainty_score":0.016626835},"labels":[],"label_agreement":null},{"id":"W2130998770","doi":"10.1145/1596517.1596518","title":"Extrinsic summarization evaluation","year":2009,"lang":"en","type":"article","venue":"ACM Transactions on Speech and Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Sixth Framework Programme","keywords":"Automatic summarization; Computer science; Task (project management); Information retrieval; Audit; Human–computer interaction; Natural language processing; Artificial intelligence","score_opus":0.024218930316466627,"score_gpt":0.29291025333555487,"score_spread":0.26869132301908827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130998770","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28006172,0.005843942,0.6303888,0.0012019642,0.0014384667,0.0037494241,0.013576097,0.02590509,0.03783436],"genre_scores_gemma":[0.66572064,0.0013493291,0.2719752,0.0006516043,0.000780337,0.0020881465,0.034933437,0.0029828935,0.019518495],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97795564,0.011013883,0.002280524,0.0018795349,0.0064198915,0.000450543],"domain_scores_gemma":[0.9409888,0.027940696,0.0033251138,0.005820352,0.021047236,0.0008779025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012293993,0.0021904244,0.0012839051,0.0030159126,0.00083377556,0.00304517,0.0016740845,0.0014544902,0.010098683],"category_scores_gemma":[0.06587509,0.0002529407,0.00080758554,0.0016037388,0.0005683884,0.0030679258,0.0018361142,0.0011024199,0.0041175596],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037885576,0.0008779663,0.009817999,0.003564328,0.000605443,0.00043589008,0.0009977736,0.014730721,0.06241746,0.0032411434,0.039600782,0.859922],"study_design_scores_gemma":[0.0010370007,0.009488849,0.0650547,0.0007347289,0.0016352694,0.0030045542,0.0029590465,0.42778963,0.289561,0.009641589,0.18836515,0.0007285698],"about_ca_topic_score_codex":0.00080244784,"about_ca_topic_score_gemma":0.0011147424,"teacher_disagreement_score":0.012293993,"about_ca_system_score_codex":0.00080430816,"about_ca_system_score_gemma":0.0008347191,"threshold_uncertainty_score":0.06501764},"labels":[],"label_agreement":null},{"id":"W2131876387","doi":"10.1145/2661829.2661935","title":"A Latent Semantic Model with Convolutional-Pooling Structure for Information Retrieval","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":687,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Explicit semantic analysis; Artificial intelligence; Natural language processing; Sentence; Word (group theory); Information retrieval; Context (archaeology); Probabilistic latent semantic analysis; Ranking (information retrieval); SemEval; Pooling; Semantic computing; Task (project management); Semantic Web; Semantic technology","score_opus":0.013073929773580288,"score_gpt":0.2120307665124632,"score_spread":0.1989568367388829,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131876387","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024234597,0.001219874,0.9701304,0.0006011374,0.0000638267,0.000054552496,0.0006239923,0.0014940316,0.0015775274],"genre_scores_gemma":[0.79262537,0.0014241508,0.19448021,0.00040691753,0.00015398959,0.00026821505,0.002352222,0.00014539094,0.008143495],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995865,0.000109218025,0.000027270155,0.00012346444,0.000089501904,0.00006403626],"domain_scores_gemma":[0.99957937,0.00017691412,0.00006668597,0.00007660166,0.00007598135,0.000024370009],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00094796217,0.0009199257,0.00092595327,0.0012209847,0.0003034842,0.0009733903,0.0017516655,0.0010186252,0.0018875613],"category_scores_gemma":[0.0021986638,0.00041825307,0.0012521573,0.0016967495,0.0006762594,0.0034187539,0.0007806181,0.0012255452,0.00084992003],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005878149,0.00040952646,0.0035553733,0.0004632919,0.00046753156,0.00031397326,0.00031500947,0.51910865,0.017799715,0.12010405,0.0134797385,0.3233953],"study_design_scores_gemma":[0.0000104729015,0.000030332016,0.00020071816,0.000007095259,0.000028851518,0.000024851079,0.0000069475745,0.97810215,0.0009582904,0.019945906,0.00067437114,0.000009912858],"about_ca_topic_score_codex":0.011111241,"about_ca_topic_score_gemma":0.011833321,"teacher_disagreement_score":0.011111241,"about_ca_system_score_codex":0.0016665538,"about_ca_system_score_gemma":0.0013016689,"threshold_uncertainty_score":0.022093117},"labels":[],"label_agreement":null},{"id":"W2132052677","doi":"10.14569/ijacsa.2015.060121","title":"A Survey of Topic Modeling in Text Mining","year":2015,"lang":"en","type":"article","venue":"International Journal of Advanced Computer Science and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":368,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Latent Dirichlet allocation; Topic model; Computer science; Probabilistic latent semantic analysis; Natural language processing; Artificial intelligence; Latent semantic analysis; Field (mathematics); Probabilistic logic; Information retrieval","score_opus":0.0669373198898449,"score_gpt":0.32942733013116254,"score_spread":0.26249001024131763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132052677","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006272367,0.182315,0.7961861,0.004210368,0.00084895967,0.00044488686,0.002055285,0.0013360622,0.006331056],"genre_scores_gemma":[0.12289982,0.3139892,0.53741,0.0018265186,0.0074459156,0.0017648877,0.008256959,0.0006603642,0.005746305],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99352914,0.0025767023,0.0008018898,0.001331183,0.0015732802,0.00018770443],"domain_scores_gemma":[0.9888678,0.008538554,0.00048371276,0.00079897215,0.0011462006,0.00016484798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007165624,0.0020886927,0.0037459226,0.009564827,0.0013747475,0.004509262,0.0030411787,0.0022031236,0.0027712437],"category_scores_gemma":[0.016814025,0.0011053433,0.003300958,0.018210845,0.0009602183,0.008731122,0.001771802,0.002651416,0.0025894532],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013156072,0.00025915864,0.0089750765,0.0048180474,0.0006170789,0.00027991095,0.0007470395,0.022362202,0.0015810754,0.0481176,0.029420968,0.88269037],"study_design_scores_gemma":[0.000094922805,0.00026078307,0.01068683,0.0026162604,0.00068390556,0.0020310397,0.0013936505,0.4482298,0.0040805787,0.23841876,0.2912305,0.0002729779],"about_ca_topic_score_codex":0.0040977313,"about_ca_topic_score_gemma":0.0034253562,"teacher_disagreement_score":0.009564827,"about_ca_system_score_codex":0.001565168,"about_ca_system_score_gemma":0.0023752542,"threshold_uncertainty_score":0.037895918},"labels":[],"label_agreement":null},{"id":"W2132551194","doi":"10.1109/ccece.2011.6030578","title":"Unsupervised language model adaptation using n-gram weighting","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut National de la Recherche Scientifique","funders":"","keywords":"Latent Dirichlet allocation; Perplexity; Language model; Computer science; Weighting; Topic model; n-gram; Artificial intelligence; Cluster analysis; Vocabulary; Document clustering; Natural language processing; Word (group theory); Mathematics; Linguistics","score_opus":0.1257445697567923,"score_gpt":0.26532571009815514,"score_spread":0.13958114034136285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132551194","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008692216,0.00018870246,0.9890627,0.000051463565,0.00006139501,0.00006466407,0.000055019176,0.0012010242,0.00062267145],"genre_scores_gemma":[0.30133244,0.0006928996,0.6882727,0.00028806896,0.00020551354,0.0006072218,0.0014462091,0.0011511463,0.0060038134],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99829584,0.0007885583,0.000074493866,0.00041181425,0.0003426927,0.000086582026],"domain_scores_gemma":[0.9984999,0.0007381947,0.00009557163,0.0002975902,0.00032509957,0.000043717842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017536794,0.0009901668,0.0009100917,0.0009156844,0.0004609731,0.00082435843,0.001148413,0.0006816832,0.0014284682],"category_scores_gemma":[0.0053732153,0.00038650088,0.0012021626,0.0010328504,0.0003962619,0.0017681313,0.0012644641,0.0015453642,0.0017681977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036280815,0.0004049367,0.002309338,0.0002225518,0.00041202453,0.0002394914,0.00052318117,0.16657068,0.06555813,0.014387238,0.0059752814,0.74303436],"study_design_scores_gemma":[0.000014584577,0.000057180976,0.0008136715,0.000011932692,0.000049400525,0.00012238785,0.00004508538,0.9750767,0.010225315,0.00911876,0.004424237,0.000040768045],"about_ca_topic_score_codex":0.0026914813,"about_ca_topic_score_gemma":0.003846187,"teacher_disagreement_score":0.0026914813,"about_ca_system_score_codex":0.00040602524,"about_ca_system_score_gemma":0.00067878276,"threshold_uncertainty_score":0.009274423},"labels":[],"label_agreement":null},{"id":"W2133823287","doi":"10.1016/j.jbi.2010.07.003","title":"Unsupervised grammar induction and similarity retrieval in medical language processing using the Deterministic Dynamic Associative Memory (DDAM) model","year":2010,"lang":"en","type":"review","venue":"Journal of Biomedical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Conestoga College","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Representation (politics); Grammar; Similarity (geometry); Selection (genetic algorithm); Associative property; Content-addressable memory; Unsupervised learning; Information processing; Machine learning; Information retrieval; Artificial neural network","score_opus":0.05273643276839589,"score_gpt":0.36504667130353124,"score_spread":0.31231023853513534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133823287","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021597806,0.6096333,0.35913035,0.0017728754,0.00053983706,0.0001466673,0.00030274232,0.0010732957,0.0058032027],"genre_scores_gemma":[0.21633638,0.5093163,0.2626294,0.00077423075,0.0011492905,0.00029662534,0.0014601843,0.0001649044,0.007872673],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99968946,0.0000760863,0.000038950126,0.000079310776,0.000098729935,0.000017562092],"domain_scores_gemma":[0.9991186,0.00057185505,0.00005508854,0.00007927426,0.00015960007,0.000015472133],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011014222,0.000536653,0.0014644739,0.0010902907,0.00020007967,0.00092615007,0.0019865155,0.00081479974,0.0012455324],"category_scores_gemma":[0.0017978494,0.00034056048,0.00071893254,0.0017360151,0.00070633524,0.0020163215,0.00069270504,0.0009629146,0.00091846084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000060696493,0.00010274314,0.00055206823,0.001707852,0.00013085698,0.00008406125,0.00005646313,0.009031694,0.0030999586,0.009878107,0.003938914,0.9713567],"study_design_scores_gemma":[0.00018862548,0.00065772876,0.0076846285,0.0014730056,0.00084082596,0.0037915772,0.0003444294,0.5182005,0.03499131,0.22442377,0.20706563,0.00033798243],"about_ca_topic_score_codex":0.0019464607,"about_ca_topic_score_gemma":0.0022089267,"teacher_disagreement_score":0.0019865155,"about_ca_system_score_codex":0.00054632226,"about_ca_system_score_gemma":0.0012568353,"threshold_uncertainty_score":0.0058249235},"labels":[],"label_agreement":null},{"id":"W2134506009","doi":"10.1145/2484028.2484174","title":"Boosting novelty for biomedical information retrieval through probabilistic latent semantic analysis","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"U.S. National Library of Medicine; University of Illinois at Chicago; University of Wisconsin-Madison","keywords":"Novelty; Computer science; Probabilistic logic; Learning to rank; Information retrieval; Boosting (machine learning); Probabilistic latent semantic analysis; Novelty detection; Ranking (information retrieval); Baseline (sea); Rank (graph theory); Latent semantic analysis; Artificial intelligence; Divergence-from-randomness model; Machine learning","score_opus":0.028830473732393228,"score_gpt":0.25998215388328144,"score_spread":0.23115168015088822,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134506009","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.084325075,0.0052899034,0.90556526,0.00067196385,0.0001383437,0.00017311261,0.0003938141,0.002077513,0.001364906],"genre_scores_gemma":[0.72273237,0.0020831758,0.269674,0.00043094458,0.0007050919,0.00022238719,0.0017288093,0.00017325148,0.0022499852],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981426,0.0006147076,0.00014565066,0.0003159455,0.00064845424,0.00013266185],"domain_scores_gemma":[0.9954035,0.0028106058,0.0004630155,0.00040616927,0.0007902541,0.0001264327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034209758,0.0011044414,0.001879079,0.0038457282,0.0007152185,0.0013460091,0.0014900947,0.0013311945,0.001119238],"category_scores_gemma":[0.00841585,0.0004562013,0.001986926,0.003116377,0.0008172355,0.0036196848,0.0014113118,0.0015855309,0.0008651744],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010061197,0.00077875535,0.013813113,0.0007855214,0.00057436276,0.00021205487,0.0002958696,0.16374965,0.03437742,0.012978947,0.0078993235,0.7635288],"study_design_scores_gemma":[0.000079163765,0.0003607481,0.0027818677,0.000020910002,0.00014467776,0.00016716702,0.000039300176,0.9714263,0.0053211134,0.01783327,0.0017599191,0.00006562663],"about_ca_topic_score_codex":0.0022179235,"about_ca_topic_score_gemma":0.0029165708,"teacher_disagreement_score":0.0038457282,"about_ca_system_score_codex":0.00097435014,"about_ca_system_score_gemma":0.0008590598,"threshold_uncertainty_score":0.018092096},"labels":[],"label_agreement":null},{"id":"W2137006453","doi":"10.1162/coli_a_00206","title":"Towards Topic-to-Question Generation","year":2015,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Latent Dirichlet allocation; Predicate (mathematical logic); Natural language processing; Correctness; Artificial intelligence; Topic model; Similarity (geometry); Subsequence; Context (archaeology); String (physics); Information retrieval; Algorithm; Mathematics; Programming language","score_opus":0.08126377489865644,"score_gpt":0.32306504233377414,"score_spread":0.2418012674351177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2137006453","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013726847,0.0006225101,0.97820115,0.0005552895,0.00015400506,0.00036535464,0.0006025806,0.0039046234,0.001867704],"genre_scores_gemma":[0.20732924,0.00053255295,0.7800468,0.00041189528,0.0003195275,0.00088875426,0.0055727568,0.0007255433,0.004172885],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949032,0.002966461,0.00023913918,0.0008930683,0.0007591859,0.00023897048],"domain_scores_gemma":[0.9875201,0.0093567325,0.00035946988,0.0012229141,0.001302,0.00023868658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005216217,0.0014438295,0.0011515679,0.0024283633,0.0008650417,0.001728073,0.0023081836,0.002256634,0.0054711658],"category_scores_gemma":[0.018453216,0.0006162217,0.0019091634,0.0018433856,0.0007617014,0.0033137358,0.0028267521,0.002180405,0.0028086985],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082407496,0.0006193759,0.0044912905,0.0012334565,0.00023869153,0.0006740702,0.002780256,0.06062264,0.027433349,0.067514606,0.037859365,0.7957088],"study_design_scores_gemma":[0.00016225921,0.00011162724,0.0006799815,0.000057452122,0.000082311904,0.00039956975,0.00035705933,0.8532419,0.014239526,0.11093435,0.019689227,0.000044702],"about_ca_topic_score_codex":0.0014056124,"about_ca_topic_score_gemma":0.001634826,"teacher_disagreement_score":0.0054711658,"about_ca_system_score_codex":0.00089017337,"about_ca_system_score_gemma":0.0015182573,"threshold_uncertainty_score":0.027586281},"labels":[],"label_agreement":null},{"id":"W2138211541","doi":"10.1007/978-3-642-31178-9_10","title":"Learning Good Decompositions of Complex Questions","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Support vector machine; Artificial intelligence; Training set; Classifier (UML); Set (abstract data type); Simple (philosophy); Machine learning; Transformation (genetics); Natural language processing; Data mining; Information retrieval","score_opus":0.02919150543798983,"score_gpt":0.27164587961385434,"score_spread":0.24245437417586452,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138211541","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022259625,0.00072486827,0.965977,0.0010273673,0.000074939555,0.00005609872,0.0003659764,0.00086382957,0.008650397],"genre_scores_gemma":[0.2624151,0.001072285,0.71769816,0.0004210461,0.00020363215,0.0002877675,0.0026279858,0.0005709243,0.014703199],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985721,0.00065589166,0.00009241765,0.00035576854,0.00023828517,0.00008555717],"domain_scores_gemma":[0.9935777,0.004462581,0.00023365568,0.0009953524,0.0004368586,0.00029385183],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022348661,0.0010505464,0.0008839732,0.0011251116,0.0005746114,0.0019663516,0.0011012964,0.0007748983,0.013920034],"category_scores_gemma":[0.014848213,0.0009341674,0.0016460858,0.0009750401,0.0014503662,0.00664582,0.002897511,0.0038480726,0.0030204856],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026184547,0.00014775834,0.0025145295,0.00056551077,0.00015417968,0.000077276105,0.0013707278,0.02932434,0.005257464,0.47594544,0.029073218,0.45530763],"study_design_scores_gemma":[0.000037927202,0.000039863527,0.0004754128,0.000086657245,0.000037030994,0.00004853536,0.00020354213,0.12624638,0.0012277357,0.86079544,0.010786284,0.000015305908],"about_ca_topic_score_codex":0.0006144251,"about_ca_topic_score_gemma":0.0017490201,"teacher_disagreement_score":0.013920034,"about_ca_system_score_codex":0.0009561418,"about_ca_system_score_gemma":0.000764083,"threshold_uncertainty_score":0.046567142},"labels":[],"label_agreement":null},{"id":"W2138806976","doi":"10.1111/j.1756-8765.2010.01109.x","title":"Discovering Binary Codes for Documents by Learning Deep Generative Models","year":2010,"lang":"en","type":"article","venue":"Topics in Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":99,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Generative model; Generative grammar; Artificial intelligence; Word (group theory); Binary number; Inference; Set (abstract data type); Code (set theory); Natural language processing; Filter (signal processing); Binary code; Associative property; Deep learning; Information retrieval; Machine learning; Mathematics; Arithmetic","score_opus":0.037740180672845305,"score_gpt":0.32495399548843823,"score_spread":0.28721381481559294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138806976","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032203197,0.00042740157,0.96337664,0.00046952718,0.000050532584,0.000076219236,0.00068484526,0.0016975342,0.0010141738],"genre_scores_gemma":[0.5034541,0.00094262103,0.4855276,0.0003752471,0.00021932078,0.0003205963,0.004035402,0.00047356487,0.0046515367],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99889874,0.00030926333,0.00008059372,0.00030950943,0.00028339925,0.00011839779],"domain_scores_gemma":[0.9941451,0.0041066483,0.00046616612,0.00065369485,0.00048097532,0.00014747566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015206533,0.0010466918,0.0011215279,0.004230149,0.0006134999,0.0021334328,0.0021710435,0.0015858329,0.0028495507],"category_scores_gemma":[0.012407829,0.0009714679,0.0016443885,0.0033409868,0.001172584,0.004739951,0.0017251371,0.0027101897,0.0017796863],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045538848,0.00033891277,0.011188073,0.00046190302,0.00022769545,0.00045949093,0.00078750815,0.24564788,0.012315543,0.09721297,0.0146890795,0.61621565],"study_design_scores_gemma":[0.000020053942,0.000022195278,0.00047053728,0.00002601199,0.000026550493,0.000114019815,0.000042557567,0.9198707,0.0023829984,0.07565599,0.0013432205,0.00002519198],"about_ca_topic_score_codex":0.0056647114,"about_ca_topic_score_gemma":0.007615492,"teacher_disagreement_score":0.0056647114,"about_ca_system_score_codex":0.0017607906,"about_ca_system_score_gemma":0.0011309868,"threshold_uncertainty_score":0.01277554},"labels":[],"label_agreement":null},{"id":"W2139061136","doi":"10.1007/11863878_56","title":"MedSearch: A Retrieval System for Medical Information Based on Semantic Similarity","year":2006,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Information retrieval; Similarity (geometry); National library; Computer science; Semantic similarity; Medical information; World Wide Web; Artificial intelligence; Library science","score_opus":0.025050280224039736,"score_gpt":0.2595427266373837,"score_spread":0.23449244641334396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139061136","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06086662,0.021317754,0.7322534,0.0024179735,0.001095746,0.0015073471,0.03226345,0.12364104,0.024636727],"genre_scores_gemma":[0.16464898,0.008270001,0.7338249,0.002231741,0.0014649479,0.0012980748,0.057955008,0.004479198,0.02582712],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994455,0.000121223864,0.00008541521,0.000084233834,0.00022499898,0.000038670358],"domain_scores_gemma":[0.9990239,0.00047386557,0.000113390495,0.00009915461,0.0001822357,0.000107451626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012626118,0.0011299353,0.001677419,0.007118975,0.00046946,0.0014715932,0.0010226171,0.0013865299,0.020123437],"category_scores_gemma":[0.0038967985,0.00047337465,0.00088871346,0.0037904398,0.00034004287,0.0030179652,0.0020420484,0.00058817305,0.009971352],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018511151,0.00030904645,0.0031812708,0.0024296115,0.0004916561,0.0005908355,0.00040514767,0.0023504447,0.04638469,0.007212909,0.21924749,0.71554583],"study_design_scores_gemma":[0.0025127276,0.0021155376,0.01952568,0.0008427693,0.002520529,0.011907998,0.00097588566,0.1794783,0.12598279,0.056297354,0.5972515,0.0005889726],"about_ca_topic_score_codex":0.0009985123,"about_ca_topic_score_gemma":0.0014969099,"teacher_disagreement_score":0.020123437,"about_ca_system_score_codex":0.00050334877,"about_ca_system_score_gemma":0.00070428016,"threshold_uncertainty_score":0.06731963},"labels":[],"label_agreement":null},{"id":"W2140813566","doi":"10.1177/1557988314528238","title":"Repackaging Prostate Cancer Support Group Research Findings","year":2014,"lang":"en","type":"article","venue":"American Journal of Men s Health","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Institute of Gender and Health; Canadian Institutes of Health Research","keywords":"Psychosocial; Focus group; Formative assessment; Knowledge translation; Context (archaeology); Interactivity; Medical education; Test (biology); Psychology; Computer science; Knowledge management; World Wide Web; Medicine; Pedagogy","score_opus":0.04359463641037264,"score_gpt":0.38944209551507436,"score_spread":0.3458474591047017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140813566","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6540174,0.00543285,0.16610172,0.071209535,0.004118063,0.01688473,0.0023905453,0.002167593,0.07767753],"genre_scores_gemma":[0.7732442,0.004013524,0.18147704,0.014287929,0.0009873959,0.012120484,0.0015817586,0.0013937408,0.010893991],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.86214936,0.10524329,0.01047063,0.0043563014,0.015306508,0.0024739983],"domain_scores_gemma":[0.48825228,0.4226458,0.01700892,0.043090656,0.02719511,0.0018071735],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.19301723,0.0010673673,0.0012280698,0.0077454927,0.005127955,0.009093021,0.0036084224,0.0023307065,0.008999502],"category_scores_gemma":[0.3572666,0.0012288399,0.0013055401,0.006139052,0.0039919484,0.011984006,0.010815806,0.0042355997,0.002540879],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032961,0.0007406001,0.016241957,0.004807712,0.00018046614,0.0016142967,0.6074882,0.0004191078,0.0032906067,0.0064724027,0.015582501,0.34283254],"study_design_scores_gemma":[0.0003022967,0.0016597797,0.031043645,0.011331945,0.0006219424,0.0018870506,0.6597796,0.002537381,0.0157298,0.039145634,0.23556668,0.00039430516],"about_ca_topic_score_codex":0.002545979,"about_ca_topic_score_gemma":0.0068723448,"teacher_disagreement_score":0.19301723,"about_ca_system_score_codex":0.003745415,"about_ca_system_score_gemma":0.008142423,"threshold_uncertainty_score":0.9951534},"labels":[],"label_agreement":null},{"id":"W2140920152","doi":"10.1007/978-3-540-68825-9_27","title":"A Statistical Model for Topic Segmentation and Clustering","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Cluster analysis; Segmentation; Document clustering; Artificial intelligence; Statistical model; Brown clustering; Hierarchical clustering; Topic model; Bayesian probability; Natural language processing; Pattern recognition (psychology); Data mining; Fuzzy clustering; Canopy clustering algorithm","score_opus":0.0366026267238402,"score_gpt":0.2742441229591897,"score_spread":0.2376414962353495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140920152","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019459656,0.0004412885,0.99583673,0.00018931822,0.00005498687,0.00004602618,0.0002905944,0.0006152461,0.00057992147],"genre_scores_gemma":[0.13779785,0.0019096687,0.84024435,0.00034710584,0.0007115449,0.001069882,0.0044975635,0.0011150978,0.012306969],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966348,0.0013397308,0.00024486374,0.00089206337,0.0006476075,0.00024095006],"domain_scores_gemma":[0.9906703,0.0067075808,0.00042515868,0.001078736,0.00090192637,0.00021633397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053362614,0.0013421611,0.0023150248,0.00342689,0.0015242521,0.0035101287,0.00498053,0.0032321096,0.0055051944],"category_scores_gemma":[0.01682482,0.0015845983,0.0031882613,0.005929234,0.0015786919,0.005286216,0.002320199,0.0035276182,0.0043664183],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034163124,0.00019547806,0.0022767577,0.00046357658,0.0003758618,0.0002338087,0.0008905799,0.31437677,0.005269645,0.28469017,0.024597926,0.36628786],"study_design_scores_gemma":[0.00002055683,0.000023474398,0.00041072277,0.000025697926,0.000048106613,0.00010442228,0.000035275898,0.8550005,0.00064881996,0.139152,0.004498846,0.000031712432],"about_ca_topic_score_codex":0.010155239,"about_ca_topic_score_gemma":0.012791855,"teacher_disagreement_score":0.010155239,"about_ca_system_score_codex":0.002196394,"about_ca_system_score_gemma":0.0020762654,"threshold_uncertainty_score":0.02822113},"labels":[],"label_agreement":null},{"id":"W2141520705","doi":"10.1023/a:1026028229881","title":"Applying Machine Learning to Text Segmentation for Information Retrieval","year":2003,"lang":"en","type":"article","venue":"Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":76,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs","keywords":"Segmentation; Computer science; Text segmentation; Artificial intelligence; Pattern recognition (psychology); Natural language processing; Word (group theory); Mathematics","score_opus":0.016645412208762866,"score_gpt":0.25568494074736936,"score_spread":0.2390395285386065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141520705","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016460698,0.0018568216,0.9744391,0.00042954506,0.00016642034,0.00019143653,0.00023332499,0.004581825,0.0016408168],"genre_scores_gemma":[0.19853266,0.0016855876,0.7919687,0.0002758656,0.00051196327,0.00032763428,0.0017109981,0.0008884883,0.0040980643],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99805945,0.0009328028,0.0001661074,0.0003937302,0.00031570293,0.00013220888],"domain_scores_gemma":[0.994518,0.0039974777,0.0002636971,0.00048577698,0.0006177155,0.000117413045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023408048,0.0011852646,0.0016790185,0.004115161,0.0011106795,0.0025003417,0.001268672,0.0015176508,0.0031278236],"category_scores_gemma":[0.00960903,0.00071805646,0.0015518145,0.003868007,0.00088751526,0.0032651532,0.0010387112,0.0017421303,0.0029929597],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033933713,0.00036011872,0.0018411519,0.000602051,0.0002307993,0.00013779118,0.00042240272,0.05574257,0.038491827,0.0083670085,0.011189953,0.8822751],"study_design_scores_gemma":[0.000044603945,0.00011037235,0.0012125934,0.00003842956,0.000112274596,0.000107288884,0.00014638873,0.9312439,0.020282906,0.04036074,0.006301579,0.00003882279],"about_ca_topic_score_codex":0.0062461197,"about_ca_topic_score_gemma":0.006248933,"teacher_disagreement_score":0.0062461197,"about_ca_system_score_codex":0.0012341846,"about_ca_system_score_gemma":0.0014211114,"threshold_uncertainty_score":0.012419522},"labels":[],"label_agreement":null},{"id":"W2141694739","doi":"10.1145/544220.544227","title":"Using librarian techniques in automatic text summarization for information retrieval","year":2002,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; National Science Foundation","keywords":"Automatic summarization; Computer science; Information retrieval; Multi-document summarization; Operationalization; World Wide Web; Seekers","score_opus":0.04868112608135101,"score_gpt":0.2590906807895311,"score_spread":0.21040955470818012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141694739","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00853607,0.0012329996,0.9808896,0.00038485148,0.00004940771,0.00026957475,0.00028253294,0.0063279453,0.0020271025],"genre_scores_gemma":[0.059675153,0.0011400895,0.93277574,0.0001930028,0.0002158756,0.00042033143,0.0016564768,0.0006925773,0.0032307955],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99497473,0.0026750418,0.0005199079,0.00080832787,0.00088615256,0.0001359554],"domain_scores_gemma":[0.9898947,0.0060477117,0.0009540849,0.0011139293,0.001872063,0.00011758988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059250584,0.0016083254,0.0012740326,0.00933485,0.0013880448,0.0032392116,0.0015838569,0.0009806154,0.0044291937],"category_scores_gemma":[0.016989563,0.00078495574,0.0013020545,0.0077946223,0.0008452297,0.00488404,0.0016415094,0.0014907137,0.0046081655],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016635333,0.0001674992,0.0025705923,0.00080763485,0.00021246134,0.00011804589,0.002860034,0.008778193,0.023035936,0.009179954,0.00932061,0.9427827],"study_design_scores_gemma":[0.00035575422,0.0016169468,0.018761754,0.00057253725,0.0015163128,0.0016412577,0.0041803583,0.47512186,0.17747521,0.08273021,0.2353119,0.0007159044],"about_ca_topic_score_codex":0.0023854545,"about_ca_topic_score_gemma":0.0046373736,"teacher_disagreement_score":0.00933485,"about_ca_system_score_codex":0.00073529774,"about_ca_system_score_gemma":0.0011052634,"threshold_uncertainty_score":0.031335056},"labels":[],"label_agreement":null},{"id":"W2142018130","doi":"","title":"DalTREC 2006 QA System Jellyfish: Regular Expressions Mark-and-Match Approach to Question Answering","year":2006,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Question answering; Rewriting; Computer science; Jellyfish; Robustness (evolution); Regular expression; Artificial intelligence; Architecture; Natural language processing; Information retrieval; Programming language","score_opus":0.02211898084258234,"score_gpt":0.2386022074487429,"score_spread":0.21648322660616054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142018130","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021564543,0.0011419544,0.70789266,0.0021198795,0.00059181533,0.0014779385,0.008540386,0.23761956,0.019051244],"genre_scores_gemma":[0.1391872,0.00066149567,0.7768408,0.0018601711,0.00030573749,0.000841115,0.027915915,0.0066673285,0.04572025],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972149,0.00057547033,0.0002944801,0.0008810735,0.00086407276,0.00016993731],"domain_scores_gemma":[0.9959078,0.0008428835,0.00020509765,0.0012992956,0.0015284701,0.00021648462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004331584,0.0012993452,0.0016603756,0.002724305,0.0011612981,0.0026325553,0.003740552,0.0023580359,0.021742456],"category_scores_gemma":[0.00815233,0.0009049403,0.0015033173,0.0012385584,0.0011274472,0.005056226,0.003380413,0.0031409469,0.016759612],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012449847,0.0007395202,0.003054888,0.0015461256,0.0003236627,0.000505943,0.0013110933,0.014273748,0.10934965,0.029407611,0.32714584,0.51109695],"study_design_scores_gemma":[0.0004850458,0.0008989056,0.0030641719,0.00014796031,0.00026245118,0.0011487528,0.0003465562,0.37211236,0.15032706,0.035016958,0.43579692,0.00039277677],"about_ca_topic_score_codex":0.012273865,"about_ca_topic_score_gemma":0.010109714,"teacher_disagreement_score":0.021742456,"about_ca_system_score_codex":0.0021492213,"about_ca_system_score_gemma":0.0029680212,"threshold_uncertainty_score":0.07273579},"labels":[],"label_agreement":null},{"id":"W2142564378","doi":"","title":"Distributed EDLSI, BM25, and Power Norm at TREC 2008","year":2008,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Computer science; Normalization (sociology); Weighting; Information retrieval; Search engine indexing; Vector space model; Data mining; Relevance feedback; Relevance (law); Artificial intelligence; Image retrieval","score_opus":0.03705805193682006,"score_gpt":0.24251510019953243,"score_spread":0.20545704826271238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142564378","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50603604,0.014326471,0.2505009,0.008222388,0.0072657913,0.0087485425,0.05249018,0.06644563,0.08596397],"genre_scores_gemma":[0.64658755,0.0008632124,0.24010946,0.0013037252,0.0010527514,0.003308867,0.0809052,0.0019256015,0.023943648],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9804387,0.006965431,0.0011705604,0.0034551236,0.0068779257,0.0010923109],"domain_scores_gemma":[0.9806974,0.0046479544,0.0007798347,0.0048809214,0.007594686,0.001399185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02397801,0.0028022074,0.002992799,0.0045212945,0.0025705441,0.0029803808,0.0042262506,0.0025137784,0.0068952222],"category_scores_gemma":[0.03275941,0.0006501659,0.0013023484,0.003851366,0.0013574841,0.0041917614,0.0031447904,0.0039422954,0.0054177297],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004258039,0.0037506174,0.007264738,0.0014127198,0.0007542355,0.00026164495,0.0004359516,0.06414228,0.020581966,0.0041170064,0.3172338,0.575787],"study_design_scores_gemma":[0.0042232675,0.0054255193,0.036562447,0.00021545905,0.0004951401,0.00054056273,0.0008396891,0.7540265,0.10816098,0.013087308,0.07574192,0.00068111904],"about_ca_topic_score_codex":0.024245862,"about_ca_topic_score_gemma":0.03495745,"teacher_disagreement_score":0.024245862,"about_ca_system_score_codex":0.0059259715,"about_ca_system_score_gemma":0.0040288274,"threshold_uncertainty_score":0.12680936},"labels":[],"label_agreement":null},{"id":"W2143654815","doi":"10.1109/iciis.1999.810309","title":"Multi-word complex concept retrieval via lexical semantic similarity","year":2003,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"WordNet; Computer science; Semantic similarity; Natural language processing; Artificial intelligence; Word (group theory); Similarity (geometry); Taxonomy (biology); Information retrieval; Task (project management); Mathematics; Image (mathematics)","score_opus":0.07497175254081719,"score_gpt":0.2990695062509769,"score_spread":0.22409775371015972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143654815","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09424217,0.0013508401,0.89973915,0.00024125063,0.00006334767,0.00019900319,0.0002905681,0.0007911872,0.0030824586],"genre_scores_gemma":[0.628239,0.00065122754,0.36847556,0.000089034715,0.00013535307,0.00027268523,0.00063924957,0.00010066916,0.0013972291],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99794894,0.0006476413,0.00016428667,0.00044514198,0.0007152156,0.00007876676],"domain_scores_gemma":[0.9970329,0.0017317703,0.0003519916,0.000448342,0.00034259228,0.0000924124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018703876,0.0005761741,0.0013754665,0.0059725232,0.00056698325,0.0023978336,0.0012108104,0.0007780471,0.0022045947],"category_scores_gemma":[0.0107412,0.00026348137,0.0006829818,0.0055534183,0.00095431565,0.008065744,0.0024160452,0.0005619287,0.00073807046],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000651702,0.00033058948,0.0048451284,0.000840654,0.00038132098,0.00020131568,0.0010086598,0.031617273,0.03175683,0.07443538,0.004218412,0.8497127],"study_design_scores_gemma":[0.00009659388,0.00048755284,0.0065345895,0.00008657472,0.00023244353,0.00061516534,0.0006293756,0.8096518,0.017847594,0.15536311,0.008285726,0.00016944027],"about_ca_topic_score_codex":0.0011040485,"about_ca_topic_score_gemma":0.001255235,"teacher_disagreement_score":0.0059725232,"about_ca_system_score_codex":0.00071567186,"about_ca_system_score_gemma":0.00078442256,"threshold_uncertainty_score":0.009891689},"labels":[],"label_agreement":null},{"id":"W2144173701","doi":"10.3115/1667583.1667609","title":"A syntactic and lexical-based discourse segmenter","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Parsing; Computer science; Artificial intelligence; Natural language processing; Probabilistic logic; Recall; Market segmentation; Segmentation; Linguistics","score_opus":0.018125173145590166,"score_gpt":0.275855511870072,"score_spread":0.25773033872448187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144173701","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015568331,0.00029217938,0.9199413,0.00037266145,0.00014901858,0.00025110593,0.0017878753,0.057437327,0.0042001638],"genre_scores_gemma":[0.07150134,0.00024314264,0.90972924,0.0002609264,0.00016112259,0.00031949254,0.0040554837,0.0035076793,0.010221494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998514,0.0003020889,0.000115752104,0.0005510277,0.00044688486,0.00007037928],"domain_scores_gemma":[0.9944459,0.0031699531,0.00034152105,0.0007690665,0.0011030004,0.00017053515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023065754,0.0012176719,0.0011401888,0.0025129092,0.000928177,0.0017873829,0.0016969265,0.0017488635,0.01819033],"category_scores_gemma":[0.008151984,0.00091837905,0.0007528587,0.0013436787,0.000820027,0.004796744,0.0018302822,0.0020661317,0.011047668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006089143,0.00024927748,0.0026511285,0.00076899095,0.000104180144,0.00044257252,0.0021505146,0.004845155,0.16458233,0.01742267,0.037487064,0.7686872],"study_design_scores_gemma":[0.00023292482,0.0005541556,0.0058417325,0.00015825019,0.00030314154,0.0020654358,0.0011078109,0.3530687,0.37972665,0.026990078,0.22960566,0.00034540717],"about_ca_topic_score_codex":0.0013397028,"about_ca_topic_score_gemma":0.0016090309,"teacher_disagreement_score":0.01819033,"about_ca_system_score_codex":0.00050662074,"about_ca_system_score_gemma":0.0011361144,"threshold_uncertainty_score":0.060852706},"labels":[],"label_agreement":null},{"id":"W2144719390","doi":"10.1162/coli_a_00077","title":"Information Status Distinctions and Referring Expressions: An Empirical Study of References to People in News Summaries","year":2011,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Automatic summarization; Computer science; Salience (neuroscience); Affect (linguistics); Context (archaeology); Information retrieval; Natural language processing; Artificial intelligence; Linguistics; History","score_opus":0.09209636236490953,"score_gpt":0.33507287718669815,"score_spread":0.24297651482178861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144719390","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9927873,0.00037925752,0.003992212,0.00034751752,0.000011231908,0.000048065816,0.00016197911,0.000049582508,0.0022228332],"genre_scores_gemma":[0.9963242,0.00020511493,0.0022890323,0.00007654227,0.000023074637,0.000042479267,0.0003893563,0.000029769697,0.0006204222],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99103165,0.006907159,0.0005690258,0.0005279836,0.000839199,0.00012505037],"domain_scores_gemma":[0.7994827,0.1751615,0.014367292,0.0049472847,0.0051330905,0.0009082447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009096691,0.0002345626,0.00036844387,0.0016197737,0.0011911279,0.0015779585,0.000752532,0.0012207952,0.002076461],"category_scores_gemma":[0.15185654,0.00028962136,0.0002641439,0.0022907942,0.0012581684,0.0043621897,0.0012945832,0.0016253586,0.00047768396],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015224871,0.0010551634,0.543169,0.001358224,0.00029092943,0.0022198202,0.26901254,0.0035975478,0.009458583,0.0121398745,0.006337396,0.1498384],"study_design_scores_gemma":[0.00014263009,0.0016703518,0.7745522,0.0005968929,0.00036038243,0.004066068,0.11678097,0.04575997,0.007867604,0.014671055,0.03330232,0.00022953618],"about_ca_topic_score_codex":0.0017868188,"about_ca_topic_score_gemma":0.002123353,"teacher_disagreement_score":0.009096691,"about_ca_system_score_codex":0.0004787404,"about_ca_system_score_gemma":0.00027079188,"threshold_uncertainty_score":0.048108518},"labels":[],"label_agreement":null},{"id":"W2145410165","doi":"10.1145/2600428.2609534","title":"The effect of expanding relevance judgements with duplicates","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Relevance (law); Context (archaeology); Set (abstract data type); Natural language processing; Reusability; Information retrieval; Artificial intelligence; Test set; Sentence; Data mining; Programming language; Software","score_opus":0.0068882913074410764,"score_gpt":0.2347738308926272,"score_spread":0.22788553958518615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145410165","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9754572,0.0024365364,0.015727248,0.00044978666,0.00041564737,0.00040686355,0.00035679643,0.0011172385,0.0036327785],"genre_scores_gemma":[0.96602684,0.00038468183,0.029596495,0.00039622615,0.0003575729,0.00020624512,0.0010333373,0.0003395995,0.0016589577],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9742863,0.012373088,0.0027912487,0.0044077127,0.0053396407,0.00080198597],"domain_scores_gemma":[0.69184595,0.25440195,0.0090604415,0.024124935,0.017110322,0.003456339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024604084,0.0021016293,0.0020980113,0.0019443322,0.0025643664,0.0025855051,0.001821352,0.0021835163,0.0018390503],"category_scores_gemma":[0.20790792,0.0008875424,0.0012614141,0.0017052833,0.0016007685,0.0046522515,0.0030345717,0.0034264144,0.0006356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.015530484,0.0041119163,0.05660184,0.0037589488,0.002312298,0.0024947098,0.008374381,0.10931962,0.2789847,0.0017382529,0.010005242,0.5067676],"study_design_scores_gemma":[0.0022825457,0.026935754,0.23157565,0.00090147066,0.0058544665,0.0058413628,0.0067402544,0.34350336,0.33453608,0.014381416,0.02597563,0.0014720828],"about_ca_topic_score_codex":0.0052362033,"about_ca_topic_score_gemma":0.0050125644,"teacher_disagreement_score":0.024604084,"about_ca_system_score_codex":0.0011370173,"about_ca_system_score_gemma":0.0016892981,"threshold_uncertainty_score":0.1301204},"labels":[],"label_agreement":null},{"id":"W2145908128","doi":"10.48550/arxiv.1412.6448","title":"Embedding Word Similarity with Neural Machine Translation","year":2014,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Embedding; Natural language processing; Similarity (geometry); Machine translation; Artificial intelligence; Word (group theory); Word embedding; Translation (biology); Vocabulary; Example-based machine translation; German; Class (philosophy); Function (biology); Linguistics","score_opus":0.08227646160660891,"score_gpt":0.1993835748629519,"score_spread":0.117107113256343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145908128","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058594152,0.0008312275,0.9358972,0.000541203,0.000101110665,0.00006441269,0.0002585475,0.00091121875,0.0028010048],"genre_scores_gemma":[0.7704972,0.000712186,0.22029613,0.00022132965,0.00020404666,0.00027316413,0.0011049788,0.0002412017,0.0064496966],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991541,0.00043349643,0.000053805197,0.0001997234,0.00011576491,0.00004306797],"domain_scores_gemma":[0.99761546,0.0014822544,0.00022568951,0.0003808505,0.00024569817,0.000050028517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012138186,0.00083960674,0.0008062138,0.0010790623,0.00034792017,0.001141028,0.0010209107,0.0012093212,0.00254303],"category_scores_gemma":[0.010263956,0.00045486973,0.0008377507,0.0016874011,0.0008033063,0.0037179324,0.0016691524,0.0015124133,0.0009223959],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019490733,0.00013968933,0.0019597802,0.0002551711,0.00016386613,0.00013394978,0.00022208014,0.6668756,0.0028812452,0.06807817,0.003607256,0.25548837],"study_design_scores_gemma":[0.000008256129,0.000019344825,0.00013156225,0.0000080347445,0.000008357737,0.000018376777,0.000012883434,0.9446552,0.0005579368,0.053968485,0.00060428836,0.000007302718],"about_ca_topic_score_codex":0.0025498646,"about_ca_topic_score_gemma":0.0028102575,"teacher_disagreement_score":0.0025498646,"about_ca_system_score_codex":0.00088667386,"about_ca_system_score_gemma":0.00056424376,"threshold_uncertainty_score":0.008507311},"labels":[],"label_agreement":null},{"id":"W2146546639","doi":"","title":"Learning Bigrams from Unigrams","year":2008,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Bigram; Perplexity; Computer science; Language model; Natural language processing; Artificial intelligence; Word (group theory); Trigram; Oracle; Set (abstract data type); Speech recognition; Mathematics; Programming language","score_opus":0.022393896773708673,"score_gpt":0.24455504663165106,"score_spread":0.22216114985794239,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146546639","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1088765,0.0011678617,0.881627,0.0004976408,0.00015788536,0.000094216804,0.0007399554,0.004354657,0.0024841614],"genre_scores_gemma":[0.64789975,0.0009393764,0.33947104,0.00031786438,0.00023041801,0.0002246472,0.0034314722,0.0006826097,0.0068028737],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990152,0.00038465482,0.000076205724,0.00027000738,0.00014191394,0.00011190582],"domain_scores_gemma":[0.99555343,0.0029371232,0.00027390893,0.0006914452,0.00041863663,0.00012553591],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017826995,0.0012680431,0.0010788671,0.0033352089,0.000702576,0.0013675707,0.0015675573,0.0013942198,0.0030891895],"category_scores_gemma":[0.010073748,0.0007294053,0.0010378715,0.0019878782,0.0006651078,0.004932363,0.001828956,0.0023234587,0.0029764713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064482796,0.00027192305,0.0058782287,0.00050950045,0.00023047328,0.0006778223,0.0009754828,0.06455686,0.014482225,0.029298857,0.010209955,0.8722639],"study_design_scores_gemma":[0.000050266975,0.00016424165,0.0012470359,0.000082423016,0.00006265422,0.00030871233,0.00034008056,0.8785776,0.0089916615,0.10546125,0.004651956,0.000062133615],"about_ca_topic_score_codex":0.0016208283,"about_ca_topic_score_gemma":0.002692252,"teacher_disagreement_score":0.0033352089,"about_ca_system_score_codex":0.00054708414,"about_ca_system_score_gemma":0.0009781776,"threshold_uncertainty_score":0.0103343725},"labels":[],"label_agreement":null},{"id":"W2146823374","doi":"10.7939/r3-b3jb-nw13","title":"Temporal abstraction in temporal-difference networks","year":2006,"lang":"en","type":"article","venue":"University of Alberta Library","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Generalization; Abstraction; Temporal difference learning; Computer science; Variety (cybernetics); Artificial intelligence; Theoretical computer science; Sequence (biology); Machine learning; Mathematics; Reinforcement learning","score_opus":0.008525590803184394,"score_gpt":0.16739516477590016,"score_spread":0.15886957397271578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146823374","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019575268,0.00012483112,0.97780025,0.0002470141,0.000020259995,0.000027131895,0.00014565788,0.00022433931,0.0018352557],"genre_scores_gemma":[0.6534814,0.00035237163,0.34034878,0.00019145911,0.000047824895,0.00019654282,0.0005844357,0.000101428544,0.004695672],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99932384,0.00023594296,0.000048339483,0.00020025174,0.00013469189,0.000056927634],"domain_scores_gemma":[0.9968214,0.002092093,0.00023968708,0.0005110341,0.00020607629,0.00012974319],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014712597,0.00052727805,0.0005413309,0.000807342,0.00055213145,0.001226556,0.0017314795,0.0010287146,0.004330869],"category_scores_gemma":[0.008805721,0.00040411478,0.00094652764,0.0010456034,0.001353947,0.0064848643,0.0019900354,0.0021962067,0.00036976303],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015953254,0.000053341395,0.0014395249,0.000120361314,0.000043181157,0.00014916541,0.00048550154,0.35639113,0.0029243005,0.52440155,0.0012081183,0.112624295],"study_design_scores_gemma":[0.000008990738,0.000017186743,0.00012943213,0.00001133613,0.000009995269,0.000034231824,0.000024760478,0.68938375,0.0006776881,0.30797276,0.0017211743,0.000008692368],"about_ca_topic_score_codex":0.0040245745,"about_ca_topic_score_gemma":0.004616419,"teacher_disagreement_score":0.004330869,"about_ca_system_score_codex":0.0015290391,"about_ca_system_score_gemma":0.0006387669,"threshold_uncertainty_score":0.01448822},"labels":[],"label_agreement":null},{"id":"W2148488947","doi":"","title":"University of Lethbridge's Participation in TREC-2004 QA Track","year":2004,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Question answering; Computer science; Information retrieval; Theme (computing); Reading (process); World Wide Web; Open domain; Track (disk drive); Natural language; Natural language processing; Linguistics","score_opus":0.028360667126079653,"score_gpt":0.2481596903004874,"score_spread":0.21979902317440775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148488947","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04666889,0.012858149,0.027216498,0.124495916,0.038038358,0.0065954216,0.17964423,0.020080056,0.5444024],"genre_scores_gemma":[0.055531565,0.0025381504,0.02714813,0.009771189,0.0034853055,0.0018802028,0.22164384,0.0023039265,0.6756977],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99004996,0.0037728134,0.00028595657,0.0011251606,0.003709466,0.0010566746],"domain_scores_gemma":[0.9693739,0.0043787495,0.00039957085,0.0020370593,0.01694592,0.0068646944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022490501,0.0017665488,0.0025061283,0.0050854865,0.0059101167,0.0075526363,0.0027736707,0.0033536525,0.12368483],"category_scores_gemma":[0.020800149,0.0006676278,0.0005852434,0.0054083774,0.0013752411,0.0046271216,0.0044518653,0.0033917427,0.052941762],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011315409,0.00019139721,0.00022292403,0.00009631658,0.000010583489,0.00003580369,0.00020082029,0.00015588361,0.0013511049,0.00077105797,0.97921306,0.017637888],"study_design_scores_gemma":[0.00019551924,0.00013520123,0.003865699,0.000087804954,0.0000295025,0.00005674447,0.00049366645,0.0022075267,0.0034292631,0.0019066329,0.9875323,0.000060285507],"about_ca_topic_score_codex":0.113711536,"about_ca_topic_score_gemma":0.20113021,"teacher_disagreement_score":0.12368483,"about_ca_system_score_codex":0.008440773,"about_ca_system_score_gemma":0.015513759,"threshold_uncertainty_score":0.41376698},"labels":[],"label_agreement":null},{"id":"W2151069987","doi":"","title":"Automatic Answer Typing for How-Questions","year":2007,"lang":"en","type":"article","venue":"North American Chapter of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Typing; Computer science; Information retrieval; Natural language processing; World Wide Web; Artificial intelligence; Speech recognition","score_opus":0.018874629432145323,"score_gpt":0.27306325356320904,"score_spread":0.2541886241310637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151069987","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04636135,0.00047712625,0.8867611,0.0007287882,0.0002887681,0.0008169408,0.00785917,0.052235045,0.004471771],"genre_scores_gemma":[0.23963687,0.00029814508,0.71281034,0.00053685054,0.00029626393,0.001169139,0.031026958,0.006025112,0.008200328],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99090695,0.0027055675,0.0010174768,0.0024562804,0.0023581574,0.00055556186],"domain_scores_gemma":[0.95922285,0.02327768,0.0025742117,0.006744175,0.007175687,0.0010054116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053799115,0.0017335056,0.0025430322,0.0060450835,0.0013452541,0.0038955738,0.0021722398,0.001965271,0.009997705],"category_scores_gemma":[0.041231066,0.0011131692,0.0014868262,0.0032364263,0.0009100765,0.0074116453,0.0046886196,0.0027643011,0.0071642743],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006932762,0.0005919515,0.037204616,0.0018172128,0.00025409862,0.0004892935,0.006998597,0.002434586,0.059408993,0.02934752,0.05858085,0.802179],"study_design_scores_gemma":[0.000242486,0.00046974336,0.03742491,0.00072972284,0.00033839763,0.0020206724,0.005355587,0.3838231,0.16976495,0.1531832,0.24612331,0.00052391266],"about_ca_topic_score_codex":0.002288708,"about_ca_topic_score_gemma":0.0037708643,"teacher_disagreement_score":0.009997705,"about_ca_system_score_codex":0.0007323854,"about_ca_system_score_gemma":0.0016771908,"threshold_uncertainty_score":0.033445597},"labels":[],"label_agreement":null},{"id":"W2151209093","doi":"10.1007/978-3-662-44659-1_7","title":"Semantic Analysis-Enhanced Natural Language Interaction in Ubiquitous Learning","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in educational technology","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"Central South University; Natural Sciences and Engineering Research Council of Canada; Athabasca University","keywords":"Automatic summarization; Computer science; Natural language processing; Natural (archaeology); Natural language; Artificial intelligence; Key (lock); Human–computer interaction; Semantic similarity","score_opus":0.007734941276094363,"score_gpt":0.27293082643694266,"score_spread":0.2651958851608483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151209093","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020382373,0.00078325055,0.9676265,0.00015191603,0.00011045171,0.0000547656,0.00020047084,0.0045100437,0.0061801267],"genre_scores_gemma":[0.44811276,0.0008525904,0.53380704,0.00018456718,0.000083787454,0.00017081203,0.00124673,0.0008045066,0.014737305],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995017,0.00015025916,0.000033725664,0.00012334759,0.00014779242,0.00004321076],"domain_scores_gemma":[0.99958855,0.00024760465,0.000016180225,0.000058883663,0.000069115435,0.000019672514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005339567,0.0003940424,0.0005340432,0.00041321543,0.0003371713,0.0011289669,0.0007535786,0.0006494097,0.0055957832],"category_scores_gemma":[0.0011384675,0.00019582605,0.00062920037,0.00056505564,0.00038384047,0.0022678496,0.0013853182,0.00070272305,0.0016984321],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004776772,0.00035862025,0.00065739994,0.00043688636,0.00007088816,0.00027727307,0.001234765,0.01649839,0.11539764,0.06778929,0.014966914,0.7818341],"study_design_scores_gemma":[0.000057800153,0.00022616275,0.0018795894,0.0000705825,0.00009602872,0.00058522343,0.0006184072,0.7048794,0.08253258,0.14208065,0.0669019,0.00007166797],"about_ca_topic_score_codex":0.0011434463,"about_ca_topic_score_gemma":0.0015479184,"teacher_disagreement_score":0.0055957832,"about_ca_system_score_codex":0.0002982711,"about_ca_system_score_gemma":0.00041111317,"threshold_uncertainty_score":0.018719733},"labels":[],"label_agreement":null},{"id":"W2151552691","doi":"10.1162/089120102762671963","title":"Generating Indicative-Informative Summaries with SumUM","year":2002,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":140,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"Université de Montréal; McGill University","keywords":"Automatic summarization; Computer science; Identification (biology); Information retrieval; Natural language processing; Process (computing); Artificial intelligence","score_opus":0.02829016747849478,"score_gpt":0.24873866525715468,"score_spread":0.2204484977786599,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151552691","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09463151,0.0017907798,0.8724024,0.0005623887,0.000254802,0.0005361673,0.0014252147,0.023758642,0.0046381922],"genre_scores_gemma":[0.25273958,0.0006821218,0.73593503,0.00016665408,0.00019146474,0.0004197634,0.0045232074,0.00096161524,0.0043804487],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99774593,0.0011828422,0.00022472306,0.0003146676,0.00046885197,0.0000629391],"domain_scores_gemma":[0.9915404,0.0051701977,0.00077435275,0.00089412235,0.0014328478,0.00018818225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032882001,0.0010263264,0.00087217847,0.0017266942,0.00067557755,0.002016862,0.0010175034,0.0007514986,0.0039211656],"category_scores_gemma":[0.015502134,0.000294746,0.00050501304,0.0015745425,0.000339388,0.0028734638,0.0014833369,0.0006341445,0.002071163],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010147734,0.00028653452,0.002052643,0.00186472,0.00019743889,0.00040943155,0.0029200108,0.021607803,0.047725454,0.009200629,0.017266657,0.8954539],"study_design_scores_gemma":[0.00051412795,0.0019362457,0.0050258613,0.0002635838,0.0006082543,0.00085840863,0.0022844598,0.61609757,0.23319963,0.021204043,0.117760174,0.00024768774],"about_ca_topic_score_codex":0.0005752751,"about_ca_topic_score_gemma":0.0009960994,"teacher_disagreement_score":0.0039211656,"about_ca_system_score_codex":0.0003729771,"about_ca_system_score_gemma":0.00054080534,"threshold_uncertainty_score":0.017389834},"labels":[],"label_agreement":null},{"id":"W2151752770","doi":"10.1023/b:inrt.0000011209.19643.e2","title":"Augmenting Naive Bayes Classifiers with Statistical Language Models","year":2004,"lang":"en","type":"article","venue":"Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":247,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Naive Bayes classifier; Artificial intelligence; Computer science; Machine learning; Bayes error rate; Bayes classifier; Bayes' theorem; Classifier (UML); Conditional independence; Bayesian programming; Natural language processing; Pattern recognition (psychology); Bayes factor; Bayesian probability; Support vector machine","score_opus":0.013246135569706644,"score_gpt":0.23719751641258194,"score_spread":0.2239513808428753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2151752770","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022382993,0.0048596133,0.95856935,0.0012830836,0.00078910723,0.0002622874,0.0008154029,0.006360177,0.00467802],"genre_scores_gemma":[0.30742505,0.0035352076,0.6714925,0.0011902794,0.0017439966,0.0005941053,0.0047144527,0.00080938835,0.0084950635],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99401516,0.00296964,0.00042591227,0.00075318996,0.0015950232,0.0002410549],"domain_scores_gemma":[0.9773309,0.017233443,0.00047390812,0.0014046349,0.0033459791,0.00021105814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077082305,0.0019560058,0.0028207814,0.0044482276,0.0011497216,0.0031331896,0.0023705685,0.0024643699,0.0049339873],"category_scores_gemma":[0.031472478,0.0010347067,0.0020054833,0.0036352542,0.00061085395,0.007949124,0.001710059,0.0031287754,0.007152606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058743753,0.0006005458,0.0030491687,0.0005393695,0.00047378088,0.00014655326,0.00020473482,0.06320536,0.0057276585,0.006089303,0.02235456,0.89702153],"study_design_scores_gemma":[0.000088580775,0.00013295432,0.00073575904,0.00008684809,0.00032677408,0.00017731225,0.00007593527,0.9551525,0.003909549,0.032687556,0.0065668044,0.000059372174],"about_ca_topic_score_codex":0.0072441846,"about_ca_topic_score_gemma":0.010279382,"teacher_disagreement_score":0.0077082305,"about_ca_system_score_codex":0.00087294384,"about_ca_system_score_gemma":0.0016383493,"threshold_uncertainty_score":0.040765524},"labels":[],"label_agreement":null},{"id":"W2152380671","doi":"","title":"Open Information Extraction with Tree Kernels","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Relationship extraction; Computer science; Natural language processing; Relation (database); Artificial intelligence; Task (project management); Generalization; Sentence; Tree kernel; Tree (set theory); Information extraction; Noun; Set (abstract data type); Verb; Support vector machine; Data mining; Mathematics; Kernel method","score_opus":0.022266423387646634,"score_gpt":0.25276524020775254,"score_spread":0.2304988168201059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152380671","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0069014253,0.0003624267,0.98165786,0.00013296491,0.000042946474,0.00008049415,0.00071876653,0.008625454,0.0014776309],"genre_scores_gemma":[0.18341082,0.0005763354,0.80186117,0.00017548278,0.000086759916,0.00019041228,0.007005655,0.0011237105,0.005569645],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99821293,0.0002675721,0.00023547332,0.00054123567,0.00059163274,0.00015111265],"domain_scores_gemma":[0.9964251,0.0013086627,0.00037399287,0.0010289253,0.00074726884,0.0001159928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014635922,0.00102881,0.0012768302,0.0048432993,0.0007825258,0.0025776718,0.0016261054,0.0014927534,0.004446796],"category_scores_gemma":[0.007856127,0.0006746591,0.0016427189,0.004521378,0.0004800702,0.0070822337,0.0031234242,0.002026088,0.0054474943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023075785,0.00023413841,0.0022348599,0.0002876279,0.00012133417,0.00036338996,0.00039603602,0.014675641,0.017376889,0.028495803,0.021127505,0.91445595],"study_design_scores_gemma":[0.000031417698,0.000064379776,0.0018336899,0.00007438026,0.00008883638,0.00054758356,0.00018283927,0.79501307,0.035940446,0.12907352,0.037075378,0.000074511176],"about_ca_topic_score_codex":0.0027554256,"about_ca_topic_score_gemma":0.0036370093,"teacher_disagreement_score":0.0048432993,"about_ca_system_score_codex":0.0008569346,"about_ca_system_score_gemma":0.0010908283,"threshold_uncertainty_score":0.014876068},"labels":[],"label_agreement":null},{"id":"W2152917816","doi":"","title":"A Multi-Dimensional Bayesian Approach to Lexical Style","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Interpretability; Computer science; Natural language processing; Artificial intelligence; Word (group theory); Style (visual arts); Representation (politics); Task (project management); Bayesian probability; Linguistics","score_opus":0.03322566169584283,"score_gpt":0.2488120219387874,"score_spread":0.21558636024294459,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2152917816","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043845708,0.0003559392,0.9905682,0.0005981643,0.00007766616,0.000055407552,0.00033646772,0.0002479316,0.0033756262],"genre_scores_gemma":[0.3625523,0.0017138416,0.61730176,0.00081132824,0.0010171408,0.0008556032,0.0017057426,0.00041475284,0.013627544],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970964,0.0014212029,0.00015165006,0.00062593544,0.0005550923,0.00014969986],"domain_scores_gemma":[0.996062,0.0021509929,0.00035111766,0.00062935514,0.0006306899,0.00017582683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035575833,0.001014542,0.0014398408,0.0035382812,0.0013060364,0.0038086222,0.0025515738,0.0016503118,0.0067890636],"category_scores_gemma":[0.014595206,0.0011199163,0.00152131,0.003307972,0.001392945,0.004476082,0.0026652338,0.0023239583,0.0026532474],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015680383,0.00022489027,0.0051643145,0.00024283334,0.00028872988,0.00021626177,0.0009079182,0.12709047,0.0035702921,0.62092197,0.012034399,0.22918114],"study_design_scores_gemma":[0.000026259582,0.000042356514,0.0011983457,0.0000500898,0.000047656056,0.00017243273,0.00005789383,0.5326955,0.0004244202,0.45687017,0.008345617,0.000069195274],"about_ca_topic_score_codex":0.0043712785,"about_ca_topic_score_gemma":0.008405308,"teacher_disagreement_score":0.0067890636,"about_ca_system_score_codex":0.0013444141,"about_ca_system_score_gemma":0.0013758084,"threshold_uncertainty_score":0.022711635},"labels":[],"label_agreement":null},{"id":"W2153470478","doi":"10.5281/zenodo.8100412","title":"Short answer assessment: Establishing links between research strands","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Data science","score_opus":0.17914991403252728,"score_gpt":0.41972868532170016,"score_spread":0.24057877128917288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153470478","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13979878,0.009139563,0.6918572,0.02283402,0.0042229667,0.012847497,0.02255004,0.0061656116,0.09058435],"genre_scores_gemma":[0.31985268,0.002574811,0.6198166,0.0021938272,0.0009153862,0.019909583,0.022842182,0.0010002262,0.010894782],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9198709,0.04519195,0.012609809,0.0069051655,0.014299446,0.001122788],"domain_scores_gemma":[0.41124874,0.36910313,0.043307178,0.024896016,0.14248605,0.008958898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.09066991,0.00135883,0.0015983551,0.031826355,0.004798898,0.011974833,0.0033191042,0.0034019237,0.023653792],"category_scores_gemma":[0.51620674,0.0010528017,0.0013158614,0.017353322,0.002632802,0.022472017,0.020693516,0.0033745558,0.007680561],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014579298,0.0003647076,0.050971497,0.011011284,0.00059190945,0.000591438,0.06978178,0.0012281175,0.0050870767,0.07056201,0.105719075,0.6826332],"study_design_scores_gemma":[0.0008744332,0.0007513533,0.0434998,0.009165724,0.0012174774,0.000627362,0.06352324,0.021585293,0.0070700664,0.43905896,0.4120008,0.00062548555],"about_ca_topic_score_codex":0.0021868495,"about_ca_topic_score_gemma":0.0037464988,"teacher_disagreement_score":0.09066991,"about_ca_system_score_codex":0.0029474979,"about_ca_system_score_gemma":0.008112801,"threshold_uncertainty_score":0.47951406},"labels":[],"label_agreement":null},{"id":"W2155925704","doi":"10.3115/1609067.1609141","title":"Flexible answer typing with discriminative preference ranking","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Discriminative model; Ranking (information retrieval); Preference; Computer science; Artificial intelligence; Typing; Information retrieval; Natural language processing; Machine learning; Statistics; Mathematics; Speech recognition","score_opus":0.052043386365155125,"score_gpt":0.2595056353299653,"score_spread":0.20746224896481016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155925704","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03681551,0.0005079065,0.95495737,0.0002330625,0.00006943692,0.00019225337,0.00041043517,0.004863408,0.0019507515],"genre_scores_gemma":[0.5756209,0.00026975665,0.41596973,0.0003655452,0.00021264251,0.00032513763,0.0021865382,0.0003808275,0.0046689226],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943645,0.0027781404,0.00031856276,0.0010763954,0.0010906983,0.0003718484],"domain_scores_gemma":[0.9879957,0.0063532246,0.0007704646,0.0024116626,0.0020620914,0.00040689515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044022063,0.0013302744,0.0019237604,0.002490695,0.000628671,0.0015979665,0.0020923344,0.0017390779,0.0038453666],"category_scores_gemma":[0.01729913,0.00047242397,0.0010729469,0.0025163083,0.0005620597,0.0035502648,0.0015857348,0.002093715,0.0032903606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077152153,0.00054743595,0.0054690056,0.0003748477,0.00013339943,0.00017263267,0.00031778088,0.031478286,0.025355736,0.0067955633,0.0123240305,0.9162597],"study_design_scores_gemma":[0.000087885965,0.00034905403,0.002864997,0.000033260985,0.00006792189,0.0004494038,0.00014465369,0.94863516,0.015715575,0.026819387,0.004720873,0.0001117876],"about_ca_topic_score_codex":0.002254169,"about_ca_topic_score_gemma":0.0046591237,"teacher_disagreement_score":0.0044022063,"about_ca_system_score_codex":0.0005844187,"about_ca_system_score_gemma":0.0010048982,"threshold_uncertainty_score":0.023281336},"labels":[],"label_agreement":null},{"id":"W2156229087","doi":"10.1109/icdm.2009.107","title":"Multi-document Summarization by Information Distance","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Multi-document summarization; Information retrieval; Set (abstract data type); The Internet; Document clustering; Cluster (spacecraft); Data mining; Artificial intelligence; Cluster analysis; World Wide Web","score_opus":0.00811618052685527,"score_gpt":0.22714882137730827,"score_spread":0.219032640850453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156229087","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012982976,0.0038461727,0.97686666,0.00023198818,0.00015552374,0.0001922914,0.0006151547,0.0035839044,0.0015253964],"genre_scores_gemma":[0.1701876,0.0023071477,0.8154751,0.00011228714,0.000496531,0.00032932946,0.0046655987,0.0005903854,0.005835996],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99766386,0.0005813959,0.00025670856,0.0006291154,0.00077427697,0.000094739815],"domain_scores_gemma":[0.996711,0.0011987354,0.00048974185,0.0005598311,0.0009461238,0.00009453711],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017146188,0.0015357513,0.0019085857,0.008087994,0.000793807,0.0023661363,0.0013215049,0.00094554835,0.0022119135],"category_scores_gemma":[0.006219898,0.00050791487,0.0010465388,0.0063874647,0.0004904353,0.0036736361,0.0013105321,0.0011526797,0.0019982879],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018290032,0.0000807253,0.0008749093,0.00048217212,0.00021092955,0.00009096233,0.0002452444,0.037708636,0.0119528165,0.006379159,0.008608325,0.93318325],"study_design_scores_gemma":[0.00013736186,0.00081641873,0.0044963965,0.00012497629,0.0004903198,0.00061346754,0.00042969864,0.84249187,0.04503593,0.05630921,0.048841614,0.00021278716],"about_ca_topic_score_codex":0.001886708,"about_ca_topic_score_gemma":0.002511596,"teacher_disagreement_score":0.008087994,"about_ca_system_score_codex":0.0009325427,"about_ca_system_score_gemma":0.000809622,"threshold_uncertainty_score":0.009067893},"labels":[],"label_agreement":null},{"id":"W2157006255","doi":"","title":"A Neural Autoregressive Topic Model","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":189,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Softmax function; Computer science; Artificial intelligence; Generative model; Word (group theory); Vocabulary; Tree (set theory); Natural language processing; Autoregressive model; Representation (politics); Generative grammar; Artificial neural network; Mathematics; Linguistics; Statistics","score_opus":0.041741875460885924,"score_gpt":0.26839051556610005,"score_spread":0.22664864010521413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157006255","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011783182,0.0013687103,0.9786902,0.0010958217,0.0002117043,0.000075944445,0.0012914182,0.0013639847,0.0041190595],"genre_scores_gemma":[0.50547343,0.003799275,0.44476083,0.0011841967,0.0010314081,0.0009511912,0.006374043,0.00053486903,0.035890855],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915767,0.0002743592,0.00004357526,0.00027566028,0.00017953731,0.000069285554],"domain_scores_gemma":[0.99895096,0.00061501766,0.00008934951,0.00013489032,0.00017195358,0.00003778875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014401494,0.0008025626,0.0011093635,0.0011827241,0.0003698499,0.0017156001,0.0027125902,0.0017051274,0.003802131],"category_scores_gemma":[0.0043394514,0.00055058306,0.001208871,0.0018033534,0.00057020946,0.0031851514,0.00095040933,0.002735927,0.0027061366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029614283,0.00023436558,0.0042472724,0.00045426778,0.00037527934,0.00037968997,0.00052272144,0.49071506,0.009427182,0.1835446,0.029058034,0.2807454],"study_design_scores_gemma":[0.000015919799,0.000021838057,0.00037382753,0.000018980942,0.000034012253,0.00008272891,0.000012388059,0.9594976,0.00055670296,0.033917278,0.005452119,0.00001666733],"about_ca_topic_score_codex":0.005777945,"about_ca_topic_score_gemma":0.007348762,"teacher_disagreement_score":0.005777945,"about_ca_system_score_codex":0.00089564384,"about_ca_system_score_gemma":0.0008470704,"threshold_uncertainty_score":0.012719333},"labels":[],"label_agreement":null},{"id":"W2158139315","doi":"","title":"Word Representations: A Simple and General Method for Semi-Supervised Learning","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1948,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Chunking (psychology); Computer science; Word (group theory); Natural language processing; Artificial intelligence; Simple (philosophy); Word embedding; Speech recognition; Linguistics","score_opus":0.03624813891496347,"score_gpt":0.3391485912100758,"score_spread":0.3029004522951123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158139315","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015459596,0.00009429794,0.979254,0.00006847882,0.00008381196,0.000260254,0.0011330427,0.016649242,0.00091083103],"genre_scores_gemma":[0.03739706,0.00011578673,0.9502493,0.00013151435,0.00012012342,0.0011192556,0.004825043,0.002002099,0.0040398533],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957266,0.0013785525,0.00047703905,0.0013164532,0.0009151065,0.00018633423],"domain_scores_gemma":[0.9938969,0.0020583936,0.00028875715,0.0024082991,0.0012132472,0.000134426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004092969,0.002760967,0.0015309439,0.004293656,0.0010872564,0.0019313617,0.0037245671,0.0024117874,0.010110964],"category_scores_gemma":[0.014623295,0.0012227055,0.0020292674,0.0033687362,0.0008834621,0.0047805645,0.0032729544,0.0037511275,0.016319945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018695388,0.00024556703,0.0011946867,0.00038763802,0.00031654086,0.000081625985,0.00026859506,0.016890764,0.009251505,0.006789735,0.043433484,0.920953],"study_design_scores_gemma":[0.00016493257,0.00016596526,0.0014040583,0.00011842308,0.00016211945,0.0004137868,0.0001967962,0.8556893,0.03292648,0.065817386,0.042790316,0.00015042121],"about_ca_topic_score_codex":0.0026125743,"about_ca_topic_score_gemma":0.0057881055,"teacher_disagreement_score":0.010110964,"about_ca_system_score_codex":0.0006516128,"about_ca_system_score_gemma":0.0017465815,"threshold_uncertainty_score":0.033824563},"labels":[],"label_agreement":null},{"id":"W2158921918","doi":"","title":"York University at TREC 2006: Legal Track","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Weighting; Track (disk drive); Computer science; Term (time); Information retrieval; Text retrieval; Probabilistic logic; Order (exchange); Domain (mathematical analysis); Artificial intelligence; Data mining; Mathematics","score_opus":0.014037163529115805,"score_gpt":0.18507827711478597,"score_spread":0.17104111358567017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158921918","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053522233,0.009159159,0.020682734,0.03434412,0.021308653,0.007707844,0.4914412,0.038409643,0.3234245],"genre_scores_gemma":[0.040558062,0.0019417455,0.02726431,0.002580608,0.0014157486,0.0018506306,0.5772915,0.0022035714,0.34489384],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99581546,0.0007755524,0.00021735053,0.00047018466,0.002274316,0.00044702567],"domain_scores_gemma":[0.98080486,0.0022370992,0.0005666942,0.0019242445,0.011015007,0.003452141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011098145,0.0019383426,0.0022051707,0.005467698,0.005707731,0.0045506996,0.002163598,0.0024848352,0.08291732],"category_scores_gemma":[0.015130429,0.00073322834,0.0005767519,0.003746538,0.0008886682,0.0043910313,0.002129284,0.0034832854,0.043814294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008807661,0.00018382898,0.00035019006,0.00009381118,0.000010216057,0.000018718163,0.000037276135,0.00022938271,0.0006478085,0.00062200584,0.9780027,0.019715805],"study_design_scores_gemma":[0.00040530204,0.00030598498,0.008146981,0.00011038701,0.00005350627,0.000089564455,0.0003102687,0.008266073,0.0070132473,0.0020769965,0.97308785,0.000133759],"about_ca_topic_score_codex":0.17626522,"about_ca_topic_score_gemma":0.2990222,"teacher_disagreement_score":0.17626522,"about_ca_system_score_codex":0.0063120043,"about_ca_system_score_gemma":0.009969831,"threshold_uncertainty_score":0.35047853},"labels":[],"label_agreement":null},{"id":"W2159863202","doi":"10.1007/s10994-005-0928-7","title":"Combining Statistical Language Models via the Latent Maximum Entropy Principle","year":2005,"lang":"en","type":"article","venue":"Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs; Canada Research Chairs","keywords":"Computer science; Principle of maximum entropy; Perplexity; Language model; Artificial intelligence; Probabilistic latent semantic analysis; Probabilistic logic; Smoothing; Natural language processing; Inference; Statistical model; Entropy (arrow of time); Natural language","score_opus":0.015687245931904955,"score_gpt":0.26299922893494887,"score_spread":0.2473119830030439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2159863202","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071320585,0.0003817388,0.99091154,0.000277165,0.00004219441,0.000021359974,0.00010356653,0.0003203732,0.0008099535],"genre_scores_gemma":[0.5655903,0.0020222254,0.4216978,0.00044519838,0.00079844915,0.0004606751,0.0017121495,0.0007180562,0.0065551763],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99664956,0.0020788088,0.00013818618,0.00036840647,0.00060191576,0.00016310814],"domain_scores_gemma":[0.9903244,0.008121315,0.00031374142,0.00062248786,0.00046729328,0.00015080723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044643665,0.0011029383,0.0018811837,0.002027558,0.00068367075,0.0028687937,0.0016169994,0.00152684,0.002288891],"category_scores_gemma":[0.015299356,0.0011442351,0.0019341542,0.0017756706,0.0010099693,0.006234653,0.0031586753,0.0025071097,0.0013470472],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053870893,0.0002918291,0.0022938217,0.0004000498,0.0007520122,0.00029904436,0.0007324516,0.4750363,0.006669239,0.2338047,0.00626728,0.27291465],"study_design_scores_gemma":[0.000015318086,0.000018500466,0.00021699948,0.000012761798,0.000036495177,0.000026746317,0.0000179894,0.8866677,0.0005610676,0.111671895,0.0007309628,0.000023611445],"about_ca_topic_score_codex":0.0013417944,"about_ca_topic_score_gemma":0.0020020476,"teacher_disagreement_score":0.0044643665,"about_ca_system_score_codex":0.0006254775,"about_ca_system_score_gemma":0.0009795434,"threshold_uncertainty_score":0.023610115},"labels":[],"label_agreement":null},{"id":"W2160039544","doi":"10.3115/1220575.1220695","title":"Learning a spelling error model from search query logs","year":2005,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":109,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Spelling; Edit distance; String (physics); Artificial intelligence; Language model; Measure (data warehouse); Natural language processing; Maximization; Data mining; Mathematics; Linguistics","score_opus":0.06240637277141359,"score_gpt":0.28651131672858204,"score_spread":0.22410494395716846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160039544","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16878563,0.00052254146,0.8247105,0.00089223066,0.000089214634,0.00012393315,0.0009998505,0.0029927227,0.0008833623],"genre_scores_gemma":[0.88731956,0.0004365128,0.10531362,0.00020968454,0.00019913801,0.0002191776,0.0025609105,0.00033945404,0.0034019456],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99762565,0.00078578055,0.00018738366,0.0007823469,0.0004115435,0.00020734359],"domain_scores_gemma":[0.9815089,0.013743629,0.0012868198,0.0014835831,0.0017192493,0.00025777466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005051795,0.0013775942,0.0019266183,0.0018565395,0.00050717295,0.0018693864,0.0021508993,0.0018218134,0.0010058326],"category_scores_gemma":[0.028650597,0.0007933756,0.00096516806,0.002059401,0.0009148058,0.004835849,0.0011096367,0.0027922199,0.0010263671],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010939577,0.0006451895,0.020921668,0.00045027013,0.00028806817,0.00030113777,0.000518323,0.69519985,0.014517684,0.009166723,0.0055083428,0.2513888],"study_design_scores_gemma":[0.000017370881,0.000051080107,0.00076737715,0.0000063817965,0.00001928301,0.000043226657,0.00001882339,0.9934043,0.0019210976,0.0035194003,0.00021520382,0.000016428045],"about_ca_topic_score_codex":0.009312312,"about_ca_topic_score_gemma":0.00842378,"teacher_disagreement_score":0.009312312,"about_ca_system_score_codex":0.0013667861,"about_ca_system_score_gemma":0.0019973961,"threshold_uncertainty_score":0.026716769},"labels":[],"label_agreement":null},{"id":"W2160376432","doi":"10.3115/1629795.1629801","title":"A cognitive model for the representation and acquisition of verb selectional preferences","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Verb; Computer science; Alternation (linguistics); Argument (complex analysis); Natural language processing; Artificial intelligence; Object (grammar); Reflexive verb; Representation (politics); Set (abstract data type); Linguistics; Cognition; Modal verb; Psychology; Programming language","score_opus":0.07026216968939814,"score_gpt":0.3268636723456144,"score_spread":0.2566015026562163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160376432","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.078761466,0.0001714573,0.91058916,0.0012564643,0.00003581677,0.00012067349,0.00032915676,0.00040701556,0.008328797],"genre_scores_gemma":[0.8004723,0.00022034506,0.1953635,0.00035808704,0.00005379736,0.00039757387,0.0006451713,0.00008957249,0.0023995736],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985715,0.0005268456,0.00007202763,0.00042834654,0.00027133134,0.00012993331],"domain_scores_gemma":[0.9928752,0.0046921014,0.00059508224,0.00092543615,0.0006149449,0.00029717488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029623895,0.0005661851,0.0006118459,0.0013225399,0.00065039855,0.002872591,0.0021545526,0.0011944927,0.0046468233],"category_scores_gemma":[0.015267078,0.00071914296,0.0016863622,0.001039517,0.0020345973,0.0050247433,0.0013579392,0.0024668644,0.0005325665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048697848,0.00046194548,0.011927885,0.0004728459,0.00046750822,0.00055279647,0.004734862,0.107599355,0.021955973,0.6954881,0.0042294757,0.15162219],"study_design_scores_gemma":[0.000070293456,0.0001211597,0.0034258119,0.0000326507,0.00009148335,0.00032071883,0.0001911239,0.3781871,0.0022325562,0.6132339,0.0020319119,0.00006130268],"about_ca_topic_score_codex":0.0038149995,"about_ca_topic_score_gemma":0.003506643,"teacher_disagreement_score":0.0046468233,"about_ca_system_score_codex":0.0014874345,"about_ca_system_score_gemma":0.0011437316,"threshold_uncertainty_score":0.015666842},"labels":[],"label_agreement":null},{"id":"W2160416736","doi":"10.1145/2484028.2484098","title":"Modeling term dependencies with quantum language models for IR","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":125,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Probabilistic logic; Term (time); Computer science; Normalization (sociology); Independence (probability theory); Language model; Quantum; Matching (statistics); Representation (politics); Theoretical computer science; Artificial intelligence; Algorithm; Mathematics; Statistics","score_opus":0.03291098185111052,"score_gpt":0.24520743800695272,"score_spread":0.2122964561558422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160416736","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007112475,0.00051234994,0.98939157,0.0005136717,0.00006437257,0.00005126007,0.00021465751,0.0005050949,0.0016345833],"genre_scores_gemma":[0.5028441,0.002016759,0.48091605,0.00093280675,0.0007290542,0.00093342114,0.0015275796,0.0005437519,0.0095565105],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99769944,0.0011501735,0.00013236313,0.00039625508,0.00047077012,0.0001510133],"domain_scores_gemma":[0.9943218,0.004228491,0.00041736002,0.0005610922,0.0003684336,0.00010286063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035943284,0.0009116099,0.0014013147,0.0015782478,0.0008625355,0.002057442,0.0025468513,0.0018571182,0.0037155538],"category_scores_gemma":[0.012487399,0.0007796732,0.0017827826,0.0024276553,0.0012257536,0.0054122517,0.0017297264,0.003133063,0.0017890119],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017851772,0.00014716922,0.0010390491,0.00025278915,0.00018108252,0.00015976059,0.00032352738,0.5676045,0.0038535912,0.30845487,0.0066458345,0.111159265],"study_design_scores_gemma":[0.000010221187,0.000016640079,0.000107576,0.000008276469,0.00001840577,0.0000338491,0.000010233359,0.87927693,0.0003218912,0.11871426,0.0014628795,0.00001886953],"about_ca_topic_score_codex":0.005409342,"about_ca_topic_score_gemma":0.005434794,"teacher_disagreement_score":0.005409342,"about_ca_system_score_codex":0.0016098791,"about_ca_system_score_gemma":0.0017653434,"threshold_uncertainty_score":0.019008815},"labels":[],"label_agreement":null},{"id":"W2160825952","doi":"10.1145/1008992.1009024","title":"Dependence language model for information retrieval","year":2004,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":261,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Language model; Term (time); Smoothing; Linkage (software); Divergence-from-randomness model; Artificial intelligence; Probabilistic logic; Independence (probability theory); Graph; Sentence; Term Discrimination; Axiom; Natural language processing; Theoretical computer science; Visual Word; Mathematics; Image retrieval","score_opus":0.020587301289449723,"score_gpt":0.25835346418041455,"score_spread":0.23776616289096483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160825952","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0057595735,0.00096762885,0.9856995,0.0006978651,0.00008254485,0.00007657433,0.00065433164,0.0010436582,0.0050184033],"genre_scores_gemma":[0.54934937,0.0035691757,0.40926328,0.0014077893,0.0007845023,0.0011332615,0.0043468913,0.00083647575,0.029309254],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99809617,0.0007543518,0.00012511658,0.00036727267,0.0005280552,0.00012905449],"domain_scores_gemma":[0.99652183,0.002226819,0.00022425159,0.00046708947,0.00048189028,0.00007820795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023060169,0.00084870646,0.0011828735,0.0021802555,0.00060490187,0.0015316001,0.002291122,0.0013225192,0.004899609],"category_scores_gemma":[0.008259196,0.0005337208,0.001565557,0.0021241838,0.00083775504,0.005835143,0.0010873476,0.001996858,0.0030296217],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023747794,0.00018750799,0.0018133812,0.0003515487,0.00020479152,0.00038614558,0.0004654235,0.2084091,0.0042766132,0.62570816,0.018962706,0.13899712],"study_design_scores_gemma":[0.000017566657,0.000033363365,0.00024834427,0.000016317123,0.000037765643,0.00015727675,0.000019178546,0.81051,0.000634023,0.17840317,0.009895104,0.000027842892],"about_ca_topic_score_codex":0.0060659307,"about_ca_topic_score_gemma":0.003293014,"teacher_disagreement_score":0.0060659307,"about_ca_system_score_codex":0.0017954415,"about_ca_system_score_gemma":0.0011412319,"threshold_uncertainty_score":0.01639086},"labels":[],"label_agreement":null},{"id":"W2160938081","doi":"10.1017/s1351324912000289","title":"Modeling human newspaper readers: The Fuzzy Believer approach","year":2012,"lang":"en","type":"article","venue":"Natural Language Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Newspaper; Computer science; Fuzzy logic; Set (abstract data type); Artificial intelligence; Range (aeronautics); Natural (archaeology); Fuzzy set; Information extraction; Natural language processing; Information retrieval; Data mining; Advertising; Programming language","score_opus":0.015391496316516116,"score_gpt":0.2320307436379919,"score_spread":0.21663924732147577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160938081","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3882325,0.00036803848,0.60203904,0.0013649905,0.000027653923,0.00019641133,0.00037694164,0.0008424654,0.006551895],"genre_scores_gemma":[0.9150895,0.00009486881,0.08322735,0.000079086,0.000035959616,0.000056111072,0.00016512044,0.00002129115,0.0012307334],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835736,0.0008973602,0.0000639302,0.0003674797,0.00024497163,0.00006896384],"domain_scores_gemma":[0.9914575,0.00686009,0.00057126075,0.00040791568,0.0005336907,0.00016969092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031789837,0.0004943593,0.0004950145,0.0017747207,0.00047645663,0.0019213824,0.001207176,0.0011863862,0.0025026987],"category_scores_gemma":[0.011057295,0.00037663965,0.0006846834,0.00059246994,0.0006583535,0.0018995977,0.00066801475,0.00070998294,0.00060688436],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014422686,0.0012408026,0.111652516,0.00046734646,0.00094583625,0.0016116904,0.012025465,0.36153916,0.025820969,0.06562392,0.004946917,0.4126831],"study_design_scores_gemma":[0.000022453916,0.00012287813,0.004173066,0.00001691926,0.000069004025,0.00013032246,0.0005411913,0.9762547,0.0016532294,0.01610307,0.000884874,0.00002823764],"about_ca_topic_score_codex":0.0034821038,"about_ca_topic_score_gemma":0.0030747293,"teacher_disagreement_score":0.0034821038,"about_ca_system_score_codex":0.00064874935,"about_ca_system_score_gemma":0.00034873825,"threshold_uncertainty_score":0.016812265},"labels":[],"label_agreement":null},{"id":"W2161627020","doi":"10.1109/ictai.2008.84","title":"Answering Complex Questions Using Query-Focused Summarization Technique","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Automatic summarization; Artificial intelligence; Natural language processing; Set (abstract data type); Cosine similarity; Information retrieval; Feature (linguistics); Cluster analysis","score_opus":0.08409606663430426,"score_gpt":0.2810506007681567,"score_spread":0.19695453413385244,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161627020","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08227648,0.00068866357,0.9054908,0.00046359375,0.00005776467,0.00033773397,0.0008065137,0.0077723386,0.0021060589],"genre_scores_gemma":[0.40412283,0.0003432856,0.58880335,0.0001888268,0.0001554548,0.00024709527,0.0028962116,0.00030527808,0.002937616],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988111,0.00050078734,0.00010525006,0.0003025182,0.00022190888,0.00005843273],"domain_scores_gemma":[0.9967073,0.001963729,0.00031420155,0.0003734463,0.0005669582,0.00007433314],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001749343,0.00096629193,0.001077751,0.0010728718,0.00042117585,0.0010193591,0.0010202816,0.0010445993,0.0031867176],"category_scores_gemma":[0.0061899563,0.00025872374,0.0006189838,0.0009896377,0.00024766815,0.0018489908,0.00066089263,0.00077266153,0.0012941331],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007880031,0.000671026,0.002821447,0.00088886876,0.0003488439,0.00045028413,0.0019844328,0.05485887,0.1884333,0.0075403405,0.014588035,0.7266266],"study_design_scores_gemma":[0.00021524396,0.0010362903,0.0044335034,0.00003928975,0.0003393346,0.00064775214,0.0007960694,0.8625412,0.09596674,0.015038594,0.018827314,0.000118771786],"about_ca_topic_score_codex":0.0009789374,"about_ca_topic_score_gemma":0.001434663,"teacher_disagreement_score":0.0031867176,"about_ca_system_score_codex":0.00027344518,"about_ca_system_score_gemma":0.0003874488,"threshold_uncertainty_score":0.010660708},"labels":[],"label_agreement":null},{"id":"W2163571449","doi":"10.1007/978-3-319-10816-2_11","title":"A Topic Model Scoring Approach for Personalized QA Systems","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University; Université Laval","funders":"","keywords":"Mean reciprocal rank; Computer science; Latent Dirichlet allocation; Personalization; Similarity (geometry); Information retrieval; Question answering; Rank (graph theory); Set (abstract data type); Topic model; Probabilistic logic; Reciprocal; Ranking (information retrieval); Data mining; Artificial intelligence; World Wide Web","score_opus":0.04480884360883672,"score_gpt":0.2528119084833199,"score_spread":0.2080030648744832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163571449","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006328967,0.00070607924,0.984705,0.00031244347,0.00015734561,0.00019222045,0.0004745304,0.00443016,0.0026934403],"genre_scores_gemma":[0.20013908,0.000827761,0.7763966,0.00027036015,0.00046375292,0.00051850174,0.0038507918,0.001357614,0.016175536],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995449,0.001980131,0.00029534343,0.0006552395,0.0013273929,0.00029280642],"domain_scores_gemma":[0.9941397,0.00263196,0.00017404415,0.0010804561,0.001731703,0.00024215358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052334964,0.0012504036,0.0017011443,0.002915728,0.0017046958,0.0029768846,0.0027131373,0.002105613,0.011549768],"category_scores_gemma":[0.01630263,0.0008677974,0.0015582036,0.0035563477,0.0004971064,0.003952125,0.0027539546,0.0031187644,0.0073710396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038166074,0.00032639474,0.002057652,0.00033193838,0.00023285951,0.00013456916,0.00046973376,0.045602757,0.011233178,0.017618751,0.041744854,0.8798657],"study_design_scores_gemma":[0.000046148125,0.00012022597,0.0013239154,0.000045112352,0.00013094838,0.00021880613,0.000113274946,0.94559973,0.0063588712,0.03220162,0.013771389,0.00007002193],"about_ca_topic_score_codex":0.0052759885,"about_ca_topic_score_gemma":0.007836175,"teacher_disagreement_score":0.011549768,"about_ca_system_score_codex":0.0011331928,"about_ca_system_score_gemma":0.0019297327,"threshold_uncertainty_score":0.038637817},"labels":[],"label_agreement":null},{"id":"W2163636154","doi":"10.1007/978-3-540-85287-2_9","title":"An Efficient Statistical Approach for Automatic Organic Chemistry Summarization","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Natural language processing; Information retrieval; Artificial intelligence","score_opus":0.01678369904806987,"score_gpt":0.24083519758218888,"score_spread":0.22405149853411901,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163636154","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003785722,0.00047726455,0.98719424,0.00014361863,0.00013146811,0.00011075129,0.00076173106,0.0062504397,0.0011447507],"genre_scores_gemma":[0.05922479,0.00055924576,0.9264702,0.00016459382,0.00039374243,0.00042970452,0.005922237,0.00074142875,0.0060940417],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986733,0.00029385742,0.00013078589,0.0002531934,0.0005312339,0.00011757778],"domain_scores_gemma":[0.9978284,0.0010301593,0.00012593248,0.0002657036,0.0006861784,0.00006355503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011692145,0.0015041176,0.0017632893,0.003550322,0.0009282407,0.002005135,0.0023786556,0.0011373146,0.007847962],"category_scores_gemma":[0.0041537127,0.0007276127,0.0016046322,0.003939922,0.00043703333,0.0019670958,0.0018167893,0.0013739982,0.0064456463],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024135735,0.000144068,0.0003867421,0.00025310367,0.00012314625,0.000118111195,0.000072699055,0.031807784,0.024275813,0.008241505,0.019417048,0.91491854],"study_design_scores_gemma":[0.00006248182,0.00013987691,0.00090696756,0.000020841351,0.00010475449,0.00017053347,0.000070176284,0.94234425,0.017503269,0.023461442,0.015165263,0.000050226146],"about_ca_topic_score_codex":0.0029832274,"about_ca_topic_score_gemma":0.0061708395,"teacher_disagreement_score":0.007847962,"about_ca_system_score_codex":0.00056312385,"about_ca_system_score_gemma":0.0017490666,"threshold_uncertainty_score":0.026254058},"labels":[],"label_agreement":null},{"id":"W2164755781","doi":"10.3115/v1/d14-1085","title":"Unsupervised Sentence Enhancement for Automatic Summarization","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Grammaticality; Automatic summarization; Coreference; Computer science; Sentence; Natural language processing; Artificial intelligence; Source text; Resolution (logic); Semantics (computer science); Event (particle physics); Linguistics; Grammar; Physics","score_opus":0.020244024081394644,"score_gpt":0.2438347385067234,"score_spread":0.22359071442532877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164755781","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014557525,0.0007078592,0.97483236,0.000199989,0.00013763091,0.00021792152,0.0006132902,0.0069300015,0.0018035299],"genre_scores_gemma":[0.11832576,0.00043953987,0.8734959,0.00014757182,0.00027861958,0.0003488504,0.0031083936,0.0006716391,0.0031837462],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988005,0.0004878306,0.000105891675,0.0003108159,0.00023431338,0.000060610695],"domain_scores_gemma":[0.99655604,0.001550932,0.0003568395,0.00054389174,0.000900896,0.000091454946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018120566,0.0013422014,0.00071318954,0.0017590667,0.0005319955,0.00092729524,0.0009829332,0.00066509895,0.0044212234],"category_scores_gemma":[0.0047858213,0.00043266342,0.0008129059,0.0010682365,0.00041153588,0.0016344213,0.0010931503,0.0011725309,0.0029229717],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031377884,0.00018273221,0.0010115256,0.00075203186,0.00014401342,0.00026757465,0.0007522994,0.010231052,0.21738239,0.0061556026,0.017078085,0.7457289],"study_design_scores_gemma":[0.0001260323,0.0010100647,0.005989203,0.00016338934,0.00044292622,0.0008984319,0.0005030054,0.5724059,0.3153583,0.023394903,0.0795601,0.00014782776],"about_ca_topic_score_codex":0.00037957937,"about_ca_topic_score_gemma":0.0009773276,"teacher_disagreement_score":0.0044212234,"about_ca_system_score_codex":0.00029857722,"about_ca_system_score_gemma":0.00056792545,"threshold_uncertainty_score":0.014790475},"labels":[],"label_agreement":null},{"id":"W2166053996","doi":"10.5555/1931390.1931453","title":"Indexing low frequency information for question answering","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Question answering; Computer science; Information retrieval; Search engine indexing; Weighting; Clef; Redundancy (engineering); Feature (linguistics); Scheme (mathematics); Document retrieval; Natural language processing; Artificial intelligence; Mathematics; Linguistics","score_opus":0.012406122944990711,"score_gpt":0.2566936372010145,"score_spread":0.2442875142560238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166053996","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11551997,0.0032990796,0.85918385,0.0012861855,0.00021822873,0.0014077823,0.0017530029,0.0104336785,0.0068982327],"genre_scores_gemma":[0.28767928,0.0008260814,0.70286936,0.00034016714,0.00022117449,0.0010069394,0.0045794076,0.00034243288,0.0021352493],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9908172,0.0051021962,0.0006629691,0.0011448612,0.0019485243,0.00032425023],"domain_scores_gemma":[0.974841,0.01975728,0.00059803034,0.0024185516,0.0020648,0.00032029618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074916487,0.0009926411,0.0016053353,0.0044367807,0.0018039063,0.0028351694,0.0021098745,0.0021014032,0.0059393095],"category_scores_gemma":[0.03418128,0.00045564695,0.0012074623,0.004448251,0.0009324379,0.006823787,0.0018695573,0.0018573401,0.003533827],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014122792,0.0014286161,0.0040791053,0.001825109,0.00020810468,0.00021442006,0.002232285,0.02492958,0.058676824,0.016284779,0.01593403,0.87277484],"study_design_scores_gemma":[0.0005816713,0.0014527554,0.008025819,0.00017632548,0.0003214844,0.00091814296,0.0011669211,0.80425555,0.07628637,0.064180546,0.0423551,0.00027935766],"about_ca_topic_score_codex":0.0062692165,"about_ca_topic_score_gemma":0.005652986,"teacher_disagreement_score":0.0074916487,"about_ca_system_score_codex":0.0018186017,"about_ca_system_score_gemma":0.0013474694,"threshold_uncertainty_score":0.0396201},"labels":[],"label_agreement":null},{"id":"W2166141468","doi":"10.1093/llc/fqu061","title":"Citation segmentation from sparse &amp; noisy data: A joint inference approach with Markov logic networks","year":2014,"lang":"en","type":"article","venue":"Digital Scholarship in the Humanities","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Excellence; Scholarship; Library science; Inference; Citation; Computer science; Artificial intelligence; Philosophy; Political science; Epistemology","score_opus":0.17348657594254144,"score_gpt":0.2799996284789401,"score_spread":0.10651305253639867,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166141468","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008153373,0.00035627265,0.9893959,0.00052697037,0.000022395203,0.000035836474,0.00023094087,0.0005109399,0.00076744624],"genre_scores_gemma":[0.3596891,0.0007064658,0.63335556,0.0004621849,0.00039684164,0.0002685042,0.0017757785,0.00029197155,0.0030535886],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99610376,0.0017872477,0.00023382923,0.000929341,0.0007160929,0.00022965169],"domain_scores_gemma":[0.97517097,0.021094516,0.0014666667,0.0010599829,0.0009669828,0.00024079267],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0068884995,0.00096482557,0.0016324447,0.0051717088,0.0011720513,0.003229499,0.0030695188,0.00219044,0.0020834953],"category_scores_gemma":[0.026989834,0.0013116235,0.0022220744,0.004788768,0.00160085,0.0055138688,0.0026491496,0.0031255377,0.00063079543],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022846938,0.00018213117,0.0077840355,0.0002777117,0.0003146545,0.00041144335,0.00087975204,0.67658687,0.0022637057,0.09587008,0.0034962064,0.21170492],"study_design_scores_gemma":[0.000007238577,0.000009810883,0.0003409509,0.000013285825,0.00002117472,0.00002672496,0.000021853564,0.9503691,0.000427422,0.047871962,0.0008760701,0.00001431528],"about_ca_topic_score_codex":0.014010357,"about_ca_topic_score_gemma":0.015265177,"teacher_disagreement_score":0.9948283,"about_ca_system_score_codex":0.002836527,"about_ca_system_score_gemma":0.0022692075,"threshold_uncertainty_score":0.0364303},"labels":[],"label_agreement":null},{"id":"W2167258615","doi":"10.3115/1557690.1557694","title":"Improving the performance of the random walk model for answering complex questions","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Question answering; Natural language processing; Sentence; Artificial intelligence; Word (group theory); Graph; Information retrieval; Component (thermodynamics); Random walk; Theoretical computer science; Linguistics","score_opus":0.04762418571720622,"score_gpt":0.24355917474941643,"score_spread":0.1959349890322102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167258615","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12554172,0.001733289,0.8615733,0.0010342245,0.00019817778,0.00015973403,0.00041713665,0.0067765713,0.0025660067],"genre_scores_gemma":[0.6635083,0.000888279,0.32794863,0.0004962004,0.00028464032,0.00020281307,0.0022988976,0.0008104268,0.0035617705],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978211,0.0011857002,0.00007576715,0.00048442188,0.0002967066,0.00013635513],"domain_scores_gemma":[0.98405325,0.01356913,0.0003042525,0.00089886144,0.0009021713,0.0002722546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005007929,0.0014340564,0.0022969074,0.0014101114,0.000640169,0.0015835922,0.0017644371,0.0028702097,0.0025373357],"category_scores_gemma":[0.024175517,0.0005339857,0.00092474546,0.0012270538,0.00041864102,0.00432446,0.0009832854,0.0018848625,0.0023204905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013891001,0.000714011,0.0045087864,0.0003270315,0.0002491631,0.00014257847,0.00033739975,0.6222619,0.008764086,0.009622534,0.00918623,0.34249723],"study_design_scores_gemma":[0.000018695719,0.000038270282,0.00012951782,0.000002050238,0.00000719445,0.000008951359,0.0000074086183,0.9967615,0.0004095402,0.0024174685,0.00019293628,0.000006554808],"about_ca_topic_score_codex":0.013610858,"about_ca_topic_score_gemma":0.012906263,"teacher_disagreement_score":0.013610858,"about_ca_system_score_codex":0.0007561251,"about_ca_system_score_gemma":0.001100289,"threshold_uncertainty_score":0.02706331},"labels":[],"label_agreement":null},{"id":"W2167549729","doi":"10.1109/tfuzz.2008.2005011","title":"Domain Representation Using Possibility Theory: An Exploratory Study","year":2008,"lang":"en","type":"article","venue":"IEEE Transactions on Fuzzy Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Natural language; Representation (politics); Domain (mathematical analysis); Classifier (UML); Domain theory; Natural language understanding; Context (archaeology); Mathematics; Discrete mathematics","score_opus":0.09096865796085263,"score_gpt":0.30735610991438667,"score_spread":0.21638745195353404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167549729","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76668936,0.001063605,0.20139478,0.0015180411,0.000019788436,0.0005906358,0.00020917434,0.00015652918,0.02835809],"genre_scores_gemma":[0.9646439,0.00034124477,0.03399635,0.00006117736,0.000013432654,0.000185595,0.00009423022,0.000025664127,0.00063832494],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9902368,0.007989525,0.00022439357,0.00039107885,0.000991501,0.00016673544],"domain_scores_gemma":[0.918008,0.07723372,0.0010935398,0.0018521653,0.0015110851,0.0003013674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010846115,0.00038461166,0.0004315007,0.0025085744,0.0014773774,0.003691257,0.0011921603,0.0011461294,0.003803838],"category_scores_gemma":[0.05257673,0.00031482612,0.00053626415,0.0024808992,0.0027090902,0.008611126,0.0022266961,0.0017809853,0.00028299136],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086048903,0.0015655559,0.069466926,0.001684137,0.00015489385,0.002247229,0.19003749,0.012764549,0.010342085,0.42115545,0.0030149927,0.2867062],"study_design_scores_gemma":[0.00021207372,0.0016008699,0.06415908,0.0014395607,0.00018242311,0.0061951675,0.25693974,0.17652434,0.013435195,0.39697734,0.08206722,0.00026704726],"about_ca_topic_score_codex":0.0008443531,"about_ca_topic_score_gemma":0.0006941831,"teacher_disagreement_score":0.010846115,"about_ca_system_score_codex":0.00083312276,"about_ca_system_score_gemma":0.00061875966,"threshold_uncertainty_score":0.05736041},"labels":[],"label_agreement":null},{"id":"W2167609405","doi":"10.1613/jair.2693","title":"The Latent Relation Mapping Engine: Algorithm and Experiments","year":2008,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Princeton University","keywords":"Analogy; Relation (database); Core (optical fiber); Set (abstract data type); Latent semantic analysis; Variety (cybernetics); Relational database; Raw data","score_opus":0.2756728393136693,"score_gpt":0.4075134273255451,"score_spread":0.13184058801187576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2167609405","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6400386,0.0039928067,0.29723212,0.0020848343,0.0006197609,0.0027311041,0.003175063,0.02346424,0.026661405],"genre_scores_gemma":[0.5441714,0.00068977085,0.44427156,0.00045485646,0.000072593575,0.0017166692,0.0036584975,0.00084728404,0.004117365],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975247,0.0009866931,0.00022722928,0.0005137013,0.0005357959,0.00021183207],"domain_scores_gemma":[0.98255324,0.013557429,0.00028293114,0.0017174182,0.0015889863,0.00030001704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00528759,0.0014546667,0.0015576272,0.0011973641,0.0008447156,0.0014513964,0.003178008,0.002974618,0.0092215],"category_scores_gemma":[0.026849702,0.00062417425,0.000673834,0.0021303196,0.00072160264,0.0044035483,0.001823387,0.002008102,0.0027628278],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0045126015,0.005467864,0.009874507,0.001242687,0.00043812863,0.00035341334,0.0006001504,0.2317805,0.0064299204,0.010807319,0.031049863,0.69744307],"study_design_scores_gemma":[0.0009780178,0.00042546878,0.0014247017,0.00004688253,0.000082380175,0.00009551922,0.00018036248,0.9816638,0.0044041267,0.007859108,0.002805456,0.000034166693],"about_ca_topic_score_codex":0.01001203,"about_ca_topic_score_gemma":0.0071391347,"teacher_disagreement_score":0.01001203,"about_ca_system_score_codex":0.001328083,"about_ca_system_score_gemma":0.0020665948,"threshold_uncertainty_score":0.03084898},"labels":[],"label_agreement":null},{"id":"W2170897823","doi":"","title":"ConsentCanvas: Automatic Texturing for Improved Readability in End-User License Agreements","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Readability; Automatic summarization; USable; License; End user; World Wide Web; Information retrieval; Programming language; Operating system","score_opus":0.05487268725048904,"score_gpt":0.2629860589408521,"score_spread":0.20811337169036304,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170897823","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06821,0.0007683324,0.75802124,0.0011220933,0.00042397142,0.00074752554,0.0058761197,0.15240332,0.012427353],"genre_scores_gemma":[0.36900118,0.0006734744,0.57821774,0.00042239568,0.00032356917,0.0007525513,0.017229324,0.010691301,0.022688497],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973653,0.0012474278,0.00024549468,0.00047875533,0.00053299055,0.00012998756],"domain_scores_gemma":[0.9888876,0.0067451023,0.0009242042,0.0016304643,0.0014979949,0.00031467565],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028365608,0.0010776273,0.00063588447,0.0018786964,0.0006849004,0.0023777178,0.00121244,0.0009523025,0.019192908],"category_scores_gemma":[0.015486547,0.00051410106,0.00064844306,0.0009560985,0.0008206759,0.0060621435,0.00243667,0.0011724567,0.0071317735],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009055068,0.0002553412,0.0033178097,0.0010100881,0.00008131541,0.0006112532,0.0037980275,0.008577148,0.051109347,0.011781163,0.1338346,0.78471845],"study_design_scores_gemma":[0.000333426,0.0005673531,0.01150418,0.00041226757,0.0001460907,0.0012306916,0.002849213,0.30668637,0.24120139,0.040089462,0.394529,0.00045051955],"about_ca_topic_score_codex":0.0021369786,"about_ca_topic_score_gemma":0.0028965783,"teacher_disagreement_score":0.019192908,"about_ca_system_score_codex":0.000683307,"about_ca_system_score_gemma":0.0010115717,"threshold_uncertainty_score":0.06420666},"labels":[],"label_agreement":null},{"id":"W2171086879","doi":"10.48550/arxiv.1406.2710","title":"A Multiplicative Model for Learning Distributed Text-Based Attribute Representations","year":2014,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research","funders":"","keywords":"Computer science; Natural language processing; Word (group theory); Artificial intelligence; Variety (cybernetics); Similarity (geometry); Sentence; Context (archaeology); Multiplicative function; Linguistics; Mathematics","score_opus":0.0939435152693148,"score_gpt":0.21566468393424365,"score_spread":0.12172116866492885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171086879","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015042729,0.00014675177,0.98355603,0.0003272244,0.000042228687,0.00004635526,0.0001680548,0.00027304457,0.00039751473],"genre_scores_gemma":[0.65532166,0.00059397036,0.331293,0.00059680454,0.0003919092,0.00078437367,0.0016226907,0.00023216665,0.009163469],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982127,0.0006706416,0.00011115133,0.00059610564,0.00029075265,0.00011868455],"domain_scores_gemma":[0.99415165,0.0041874326,0.00040766931,0.0006316509,0.0004533049,0.00016826853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030452802,0.0009986764,0.0013277542,0.0014229998,0.00053947937,0.0015784317,0.0032624293,0.0017046259,0.0023671146],"category_scores_gemma":[0.012808691,0.00081855623,0.0015660983,0.0019209546,0.0013546956,0.004896183,0.001991991,0.0029646142,0.0009861684],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038212168,0.0003577843,0.0045722746,0.00029951826,0.00027607012,0.0002519748,0.0007158823,0.5824197,0.009748959,0.2140155,0.0044545582,0.18250571],"study_design_scores_gemma":[0.000016632137,0.000038282327,0.00019553772,0.000007978946,0.000020233116,0.000033474866,0.0000128212005,0.94325316,0.00067594467,0.055166706,0.0005661108,0.000013031678],"about_ca_topic_score_codex":0.0018919944,"about_ca_topic_score_gemma":0.0026765575,"teacher_disagreement_score":0.0032624293,"about_ca_system_score_codex":0.0011443709,"about_ca_system_score_gemma":0.0008237502,"threshold_uncertainty_score":0.016105175},"labels":[],"label_agreement":null},{"id":"W2173071808","doi":"10.1109/pacrim.2015.7334869","title":"New graph-based text summarization method","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Text graph; Graph; Information retrieval; Multi-document summarization; Multigraph; Natural language processing; Ranking (information retrieval); Artificial intelligence; Theoretical computer science","score_opus":0.05822045144658138,"score_gpt":0.3020652192523725,"score_spread":0.24384476780579112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2173071808","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005934675,0.0004993834,0.9879588,0.00013183456,0.00010342348,0.00012360711,0.00050155504,0.0038530943,0.00089357805],"genre_scores_gemma":[0.0986061,0.0006164909,0.8875917,0.00015726582,0.00029878013,0.0003401083,0.004205155,0.00058139185,0.007602983],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99925214,0.00012560326,0.00007384947,0.00024579637,0.00025131518,0.000051281815],"domain_scores_gemma":[0.9990245,0.00026731603,0.00012437541,0.00011029123,0.0004316767,0.000041860632],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005271776,0.0010857919,0.0010100058,0.0035367748,0.00059037725,0.0009604078,0.0011606717,0.00077176595,0.0031322218],"category_scores_gemma":[0.0020533644,0.0003538457,0.0011176527,0.002569405,0.00027839304,0.0016313111,0.00067829,0.00075078104,0.0018656576],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001796388,0.00009564366,0.0007782248,0.00048911053,0.00015976034,0.00021409022,0.0002670119,0.056766875,0.0474233,0.0065057385,0.016003292,0.8711173],"study_design_scores_gemma":[0.00007649892,0.00020489792,0.0012960777,0.000033096327,0.00018424983,0.0003373193,0.0001424393,0.93650556,0.027790051,0.008681483,0.024690827,0.000057591973],"about_ca_topic_score_codex":0.002982232,"about_ca_topic_score_gemma":0.0043535433,"teacher_disagreement_score":0.0035367748,"about_ca_system_score_codex":0.0005215547,"about_ca_system_score_gemma":0.000812206,"threshold_uncertainty_score":0.010478318},"labels":[],"label_agreement":null},{"id":"W2178198718","doi":"10.5430/air.v5n1p36","title":"The improvement of question process method in Q&amp;A system","year":2015,"lang":"en","type":"article","venue":"Artificial Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National High-tech Research and Development Program; National Natural Science Foundation of China","keywords":"Computer science; Matching (statistics); Similarity (geometry); Template; Set (abstract data type); Word (group theory); Field (mathematics); Semantic similarity; Blossom algorithm; Information retrieval; Artificial intelligence; Natural language processing; Algorithm; Data mining; Mathematics; Image (mathematics)","score_opus":0.3984937210629831,"score_gpt":0.5197412145710736,"score_spread":0.1212474935080905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2178198718","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031940702,0.00032256826,0.9580256,0.0004880836,0.00008994736,0.00086148846,0.00020180875,0.006150874,0.00191906],"genre_scores_gemma":[0.27782497,0.00023105794,0.71592087,0.00045613595,0.00010557742,0.00086629845,0.0012103651,0.00032568254,0.0030589884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97349393,0.011399882,0.002407725,0.005506475,0.006549068,0.00064298924],"domain_scores_gemma":[0.98088545,0.008204597,0.0008623443,0.0029962307,0.006457056,0.00059431716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012976414,0.0013904347,0.001281546,0.0035806887,0.0015155929,0.0028509125,0.0024829379,0.0019300983,0.004420764],"category_scores_gemma":[0.029870106,0.0005810183,0.0015960216,0.002241062,0.0011825524,0.010187889,0.002606157,0.0017821955,0.0016951462],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067369797,0.00088145724,0.013675433,0.0008732184,0.00017021745,0.00016294267,0.0026307246,0.006503563,0.03040837,0.015858494,0.0062800054,0.9218818],"study_design_scores_gemma":[0.0003833693,0.0016269728,0.018237397,0.00016253527,0.0005214379,0.0011127404,0.0016849936,0.74635834,0.12342671,0.03633694,0.069836624,0.00031183494],"about_ca_topic_score_codex":0.00442801,"about_ca_topic_score_gemma":0.0019086866,"teacher_disagreement_score":0.012976414,"about_ca_system_score_codex":0.0014914612,"about_ca_system_score_gemma":0.002831638,"threshold_uncertainty_score":0.06862664},"labels":[],"label_agreement":null},{"id":"W2179519966","doi":"10.1162/tacl_a_00104","title":"Named Entity Recognition with Bidirectional LSTM-CNNs","year":2016,"lang":"en","type":"preprint","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":115,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Feature engineering; Feature (linguistics); Word (group theory); Artificial intelligence; Lexicon; Named-entity recognition; Task (project management); Natural language processing; Encoding (memory); State (computer science); Architecture; Entity linking; Deep learning; Knowledge base","score_opus":0.02717742762268828,"score_gpt":0.2607877003144704,"score_spread":0.23361027269178214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2179519966","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07416549,0.0020486112,0.8748081,0.00087007985,0.00043674748,0.00016374876,0.0047446955,0.027812544,0.014950044],"genre_scores_gemma":[0.6360032,0.0011562776,0.32774952,0.0006165208,0.00019054915,0.00024809624,0.016910939,0.000487703,0.016637264],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994825,0.000079707104,0.000044297187,0.00020948779,0.000106627966,0.000077285134],"domain_scores_gemma":[0.9992834,0.0002051736,0.00009350174,0.00020427899,0.00018873556,0.0000250026],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073928083,0.0010313287,0.00053871935,0.0012141226,0.00035863044,0.0011497725,0.0015635727,0.0008704101,0.0037904799],"category_scores_gemma":[0.0020853232,0.00045305028,0.00057000027,0.0018834773,0.00028396837,0.0041503785,0.0013972144,0.0009036603,0.0034788426],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032681934,0.00019410552,0.002698951,0.0002794128,0.00018663848,0.00029336338,0.00013186398,0.07590875,0.029778074,0.00873275,0.029004117,0.8524651],"study_design_scores_gemma":[0.000019350608,0.00005665546,0.0011302665,0.000028789322,0.0000628322,0.00012264393,0.000053291747,0.9504978,0.02069166,0.014536359,0.012771919,0.000028440161],"about_ca_topic_score_codex":0.0072223013,"about_ca_topic_score_gemma":0.012624188,"teacher_disagreement_score":0.0072223013,"about_ca_system_score_codex":0.00088347256,"about_ca_system_score_gemma":0.00069419463,"threshold_uncertainty_score":0.014360487},"labels":[],"label_agreement":null},{"id":"W2181916342","doi":"","title":"CornPittMich Sentiment Slot-Filling System at TAC 2013","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Relevance (law); Task (project management); Sentence; Sentiment analysis; Process (computing); Knowledge base; Information retrieval; Measure (data warehouse); Base (topology); Architecture; Natural language processing; Document retrieval; Artificial intelligence; Data mining; Mathematics; Programming language","score_opus":0.008732002905667332,"score_gpt":0.21842062742979015,"score_spread":0.20968862452412282,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2181916342","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08248915,0.0016165526,0.47747472,0.0034253665,0.0012515176,0.0022605695,0.06009932,0.28134045,0.09004241],"genre_scores_gemma":[0.19460201,0.000462649,0.6195952,0.0010814565,0.00031684656,0.001456664,0.13145085,0.007348554,0.04368572],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998468,0.00031105932,0.00010276143,0.0003381573,0.0006050027,0.00017509455],"domain_scores_gemma":[0.9982626,0.00023225794,0.000064569555,0.00029403475,0.0009807441,0.00016586069],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023860706,0.0009802454,0.0010994808,0.0020405576,0.0015438653,0.0023878578,0.0019527067,0.0010761863,0.01851797],"category_scores_gemma":[0.0042429245,0.00053432677,0.00060146407,0.0017955857,0.00043393273,0.0028510413,0.0015631522,0.0016466147,0.013864959],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001217667,0.00037837977,0.0017571534,0.00029599117,0.00009828952,0.0005142455,0.0006830256,0.004297352,0.0403638,0.009073727,0.6080947,0.3332257],"study_design_scores_gemma":[0.0008560925,0.00047354627,0.0057146978,0.00014783313,0.00021486262,0.0007860277,0.0007512205,0.27338952,0.08846065,0.023433898,0.60545295,0.00031863296],"about_ca_topic_score_codex":0.024836143,"about_ca_topic_score_gemma":0.02727017,"teacher_disagreement_score":0.024836143,"about_ca_system_score_codex":0.0018186071,"about_ca_system_score_gemma":0.003402024,"threshold_uncertainty_score":0.061948776},"labels":[],"label_agreement":null},{"id":"W2182032621","doi":"","title":"Entropy-based Sentence Selection with Roget's Thesaurus","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Ottawa","funders":"","keywords":"Automatic summarization; Sentence; Computer science; Thesaurus; Information retrieval; Natural language processing; Selection (genetic algorithm); Baseline (sea); Artificial intelligence","score_opus":0.008418135404538249,"score_gpt":0.207359724092659,"score_spread":0.19894158868812073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2182032621","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28692642,0.0053311526,0.60108066,0.0031972916,0.0016950958,0.0021405602,0.009399322,0.08026146,0.009968014],"genre_scores_gemma":[0.41499105,0.0011345082,0.5413459,0.00073560595,0.0006469167,0.0009527132,0.020352684,0.002505456,0.017335212],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99644923,0.0013226544,0.0003965368,0.000728904,0.0009064119,0.00019618373],"domain_scores_gemma":[0.99278015,0.0028701383,0.00026720273,0.00084344216,0.0029800574,0.0002589694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050097755,0.001801442,0.002071983,0.0033281457,0.0012669486,0.002538613,0.0016098391,0.0015199155,0.005682188],"category_scores_gemma":[0.013082404,0.0005921758,0.001165707,0.0017606767,0.00059685414,0.0025001052,0.0018269905,0.0014039369,0.00306223],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001352878,0.00040687798,0.0027771157,0.0011347538,0.000615509,0.0004626099,0.0014869338,0.012387453,0.08498825,0.0016253808,0.04749075,0.8452716],"study_design_scores_gemma":[0.0016483717,0.004724144,0.024551893,0.0002341235,0.0017580435,0.0014267258,0.0026447156,0.501247,0.32382712,0.009173122,0.12792274,0.00084205245],"about_ca_topic_score_codex":0.012022242,"about_ca_topic_score_gemma":0.022013942,"teacher_disagreement_score":0.012022242,"about_ca_system_score_codex":0.0012904708,"about_ca_system_score_gemma":0.0021153833,"threshold_uncertainty_score":0.026494503},"labels":[],"label_agreement":null},{"id":"W2182139230","doi":"","title":"ABSUM: a Knowledge-Based Abstractive Summarizer","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Automatic summarization; Computer science; Natural language processing; Task (project management); Information retrieval; Knowledge base; Source text; Artificial intelligence; Representation (politics); Multi-document summarization; Scalability","score_opus":0.02019895973292568,"score_gpt":0.25499447608057924,"score_spread":0.23479551634765355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2182139230","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008571866,0.0013016859,0.95143926,0.00030478122,0.00023100454,0.00028884463,0.00303723,0.03131003,0.003515313],"genre_scores_gemma":[0.113538,0.0010011827,0.8585196,0.00034834995,0.00037994052,0.0006111178,0.011842198,0.0014364023,0.012323223],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918956,0.0001743031,0.00010025818,0.00018571675,0.00030123538,0.000048910228],"domain_scores_gemma":[0.99831957,0.0006194791,0.00020032312,0.00031489204,0.00046926428,0.0000765484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011021375,0.001118191,0.0010018793,0.0025475563,0.00057505415,0.001553872,0.0016941337,0.0007836904,0.006110105],"category_scores_gemma":[0.003955911,0.000447579,0.0008145824,0.0013778587,0.00030491955,0.0019013706,0.0012630456,0.0013014525,0.0049069654],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006199145,0.00012505439,0.0006602614,0.0010088446,0.00018562989,0.0002268192,0.0004980225,0.013710582,0.045056485,0.0067617847,0.037042074,0.8941046],"study_design_scores_gemma":[0.00038815103,0.0010892285,0.0044978503,0.0002352555,0.00062167813,0.00077987736,0.0005723989,0.49569497,0.15898345,0.034492824,0.3024212,0.00022321718],"about_ca_topic_score_codex":0.0016603976,"about_ca_topic_score_gemma":0.0035619354,"teacher_disagreement_score":0.006110105,"about_ca_system_score_codex":0.0005517417,"about_ca_system_score_gemma":0.0007241201,"threshold_uncertainty_score":0.02044034},"labels":[],"label_agreement":null},{"id":"W2182362006","doi":"","title":"Related Entity Finding: University of Waterloo at TREC 2010 Entity Track.","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Track (disk drive); Information retrieval; Natural language processing","score_opus":0.014957095292718397,"score_gpt":0.19904704553951932,"score_spread":0.18408995024680092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2182362006","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01702523,0.006860647,0.016331125,0.015669964,0.0029047718,0.0033559613,0.7786001,0.02173379,0.13751838],"genre_scores_gemma":[0.02516547,0.0019958543,0.021089623,0.0014549909,0.0004550975,0.00093532086,0.84229034,0.0015718519,0.10504132],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9955504,0.0007208087,0.00028328472,0.00087794976,0.0021165744,0.00045102046],"domain_scores_gemma":[0.9841001,0.0020912501,0.00052363356,0.0015076082,0.009903226,0.0018741752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006189238,0.0023502358,0.0022371686,0.008518437,0.005720084,0.0047315746,0.0038097897,0.002216997,0.07457094],"category_scores_gemma":[0.015430701,0.0010297883,0.0005942168,0.008893444,0.0011130457,0.0072461395,0.0023067256,0.002293748,0.037677817],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048486483,0.000087584514,0.000468487,0.00015647674,0.000011754887,0.000039806913,0.00008916134,0.00016681821,0.00080790016,0.00047018295,0.9820347,0.015618572],"study_design_scores_gemma":[0.0003365595,0.00010569836,0.011719392,0.00016375115,0.00006115976,0.000113537164,0.0005940302,0.0063378117,0.0066326996,0.0022771938,0.97150743,0.00015078601],"about_ca_topic_score_codex":0.63844687,"about_ca_topic_score_gemma":0.7958329,"teacher_disagreement_score":0.63844687,"about_ca_system_score_codex":0.011069963,"about_ca_system_score_gemma":0.017206002,"threshold_uncertainty_score":0.7273648},"labels":[],"label_agreement":null},{"id":"W2182378606","doi":"","title":"Universite de Montreal at TREC 2013: Experiments with Quantum Language Models in the Web Track.","year":2013,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal","funders":"","keywords":"Robustness (evolution); Computer science; Language model; Artificial intelligence; Focus (optics); Information retrieval; Data mining","score_opus":0.031722842998086316,"score_gpt":0.24765294626240222,"score_spread":0.21593010326431591,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2182378606","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63642544,0.021661274,0.06466748,0.009958145,0.006254591,0.006907433,0.1079313,0.043009583,0.10318473],"genre_scores_gemma":[0.6806932,0.0025406163,0.08554547,0.0028753835,0.0009148879,0.0026696445,0.17021467,0.0020977734,0.05244831],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99447745,0.0028328872,0.00022920218,0.0009008114,0.0010575969,0.00050210726],"domain_scores_gemma":[0.9900803,0.0054629333,0.00032787028,0.0014801621,0.0015788117,0.0010698282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011960326,0.0026300538,0.0018965852,0.001385499,0.0027202985,0.002266593,0.003127241,0.0025691898,0.011051055],"category_scores_gemma":[0.020926608,0.00077662186,0.001473482,0.0020807132,0.0009778687,0.0040662447,0.0018403193,0.003653599,0.0057027917],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005400076,0.011218584,0.009376397,0.0019056828,0.0012778353,0.00043381038,0.000850266,0.068797134,0.008614266,0.00544503,0.646936,0.23974499],"study_design_scores_gemma":[0.0052233357,0.0075320373,0.049938887,0.00040955056,0.0009850633,0.00042281713,0.0017648526,0.7205356,0.028671851,0.01795079,0.16576016,0.0008050893],"about_ca_topic_score_codex":0.291977,"about_ca_topic_score_gemma":0.38230178,"teacher_disagreement_score":0.291977,"about_ca_system_score_codex":0.0060455017,"about_ca_system_score_gemma":0.0050003836,"threshold_uncertainty_score":0.58055496},"labels":[],"label_agreement":null},{"id":"W2182794451","doi":"","title":"CUNY-BLENDER TAC-KBP2012 Entity Linking System and Slot Filling Validation System","year":2012,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Task (project management); Rank (graph theory); Cluster analysis; Data mining; Work (physics); Information retrieval; Artificial intelligence; Machine learning; Mathematics; Systems engineering; Engineering","score_opus":0.016094293216043498,"score_gpt":0.23970393194089776,"score_spread":0.22360963872485426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2182794451","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07479324,0.0015076704,0.26734498,0.0022558428,0.0012185781,0.0029635024,0.1009616,0.49621806,0.05273663],"genre_scores_gemma":[0.1826786,0.0003310035,0.38599053,0.0008561846,0.00025257556,0.0019166463,0.37983292,0.017147077,0.030994486],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99425983,0.0012544628,0.0007323126,0.001532622,0.0017690355,0.00045178537],"domain_scores_gemma":[0.9839113,0.0048331297,0.00055378285,0.003808763,0.0060359696,0.00085712905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008353065,0.002223166,0.0027062872,0.0058386875,0.0024543668,0.0040278267,0.004042345,0.002390158,0.044881724],"category_scores_gemma":[0.027074307,0.0013258802,0.001095473,0.0042517777,0.0008323792,0.010458955,0.0054120594,0.0029026212,0.026463337],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013546253,0.0006455087,0.0044874544,0.0012658858,0.00021318819,0.00085960695,0.0012963039,0.0053701648,0.01971845,0.0071097882,0.5659921,0.3916868],"study_design_scores_gemma":[0.0014421062,0.00066465494,0.011228524,0.00031620628,0.00038179316,0.0011813169,0.0016299376,0.33104363,0.0745508,0.014156139,0.56286585,0.0005389973],"about_ca_topic_score_codex":0.027823979,"about_ca_topic_score_gemma":0.028274383,"teacher_disagreement_score":0.044881724,"about_ca_system_score_codex":0.0024905957,"about_ca_system_score_gemma":0.0041301437,"threshold_uncertainty_score":0.15014434},"labels":[],"label_agreement":null},{"id":"W2183649036","doi":"","title":"RPI-BLENDER TAC-KBP2013 Knowledge Base Population System","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Knowledge base; Computer science; Graph; Base (topology); Population; Knowledge graph; Artificial intelligence; Data mining; Theoretical computer science; Mathematics","score_opus":0.011785354803987397,"score_gpt":0.23777215676114538,"score_spread":0.22598680195715798,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2183649036","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018685224,0.0006980532,0.31183085,0.000814297,0.0003582034,0.0014144917,0.15430681,0.46847412,0.043417983],"genre_scores_gemma":[0.06484703,0.0004222943,0.35516235,0.00043621496,0.00011446185,0.0012934912,0.5404471,0.018490728,0.01878633],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99770653,0.0004034396,0.00031675276,0.0006549545,0.0007881063,0.00013019246],"domain_scores_gemma":[0.9945991,0.001366309,0.000238588,0.001872505,0.0016759514,0.00024766085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004012825,0.0016343133,0.0018617378,0.006336198,0.0013745165,0.0039376975,0.004402238,0.0015003362,0.041052427],"category_scores_gemma":[0.013594396,0.001285895,0.0012609732,0.0054625324,0.00045118417,0.007214817,0.0031377233,0.0024440004,0.03729616],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012406821,0.0006809434,0.0029972,0.0013676977,0.00028332663,0.0006026947,0.0007253632,0.008109212,0.01305492,0.012667038,0.5098687,0.4484023],"study_design_scores_gemma":[0.00085236237,0.00037730334,0.0055018645,0.000300194,0.00043859647,0.001104075,0.0006718716,0.20895374,0.044532996,0.026477428,0.71049356,0.0002959517],"about_ca_topic_score_codex":0.00873426,"about_ca_topic_score_gemma":0.007430634,"teacher_disagreement_score":0.041052427,"about_ca_system_score_codex":0.0015681246,"about_ca_system_score_gemma":0.0021908183,"threshold_uncertainty_score":0.13733405},"labels":[],"label_agreement":null},{"id":"W2184384230","doi":"","title":"Transductive Learning of Structural SVMs via Prior Knowledge Constraints","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Support vector machine; Artificial intelligence; Machine learning; Labeled data; Segmentation; Pattern recognition (psychology)","score_opus":0.021911808281130377,"score_gpt":0.2701434936984218,"score_spread":0.2482316854172914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2184384230","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022215057,0.00014250503,0.97578317,0.00029662412,0.000020988025,0.000047539583,0.00007183153,0.00045271765,0.00096944004],"genre_scores_gemma":[0.7423894,0.00023637131,0.25333914,0.00034981023,0.00013758586,0.00023761734,0.00084734126,0.00014078508,0.0023219665],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986143,0.0006532519,0.00007074214,0.00031828092,0.0002662339,0.0000772376],"domain_scores_gemma":[0.99350077,0.0046122186,0.0003794823,0.0006513762,0.0007257634,0.00013033266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024101497,0.0010527807,0.0013867147,0.00083975436,0.00048141746,0.001193584,0.0020209036,0.0016089679,0.0022820996],"category_scores_gemma":[0.010387462,0.0006987103,0.00073227653,0.0009126324,0.0013897014,0.004588287,0.0017148228,0.0030026855,0.0008566011],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019949548,0.0002950262,0.0020106747,0.00019343734,0.00007980669,0.000099890785,0.00021465641,0.64754,0.0060778586,0.040448345,0.004423048,0.2984178],"study_design_scores_gemma":[0.0000046256278,0.000030379777,0.000079761616,0.0000053008816,0.0000033041106,0.000010797043,0.000007646807,0.98299825,0.00055154343,0.016131386,0.00017314126,0.0000038884627],"about_ca_topic_score_codex":0.00064845517,"about_ca_topic_score_gemma":0.001108,"teacher_disagreement_score":0.0024101497,"about_ca_system_score_codex":0.0007340407,"about_ca_system_score_gemma":0.00068878167,"threshold_uncertainty_score":0.012746215},"labels":[],"label_agreement":null},{"id":"W2184442935","doi":"","title":"JVN-TDT Entity Linking Systems at TAC-KBP2012","year":2012,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Exploit; Heuristics; Referent; Classifier (UML); Artificial intelligence; Coreference; Natural language processing; Data mining; Information retrieval; Resolution (logic)","score_opus":0.014469787648705648,"score_gpt":0.24516860369374366,"score_spread":0.230698816045038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2184442935","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12432662,0.006781884,0.35554737,0.0059312866,0.0052731195,0.0040144273,0.1313525,0.32310477,0.043668065],"genre_scores_gemma":[0.101988,0.0007547323,0.5778206,0.00092899374,0.00039762346,0.0016699106,0.29687428,0.004420369,0.015145544],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9923964,0.0021567699,0.00067923474,0.0022552244,0.0020203,0.0004920911],"domain_scores_gemma":[0.9867061,0.0040641306,0.00047734255,0.004408895,0.0036785058,0.00066513784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011690442,0.00258948,0.0021499423,0.0072749443,0.0031510347,0.0047406373,0.006535431,0.0042409184,0.011216862],"category_scores_gemma":[0.024998525,0.0013432839,0.0017148448,0.0065664104,0.00086970377,0.009095707,0.004895507,0.004222299,0.0104591865],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011885961,0.0014644047,0.0056976737,0.0012847058,0.00062360207,0.0006223734,0.0007716263,0.021712642,0.010270942,0.00618933,0.48912823,0.46104577],"study_design_scores_gemma":[0.0009709341,0.00051444845,0.00638591,0.00021938466,0.0003825881,0.00075828517,0.00088024384,0.65108,0.036473762,0.012789015,0.28925675,0.00028863625],"about_ca_topic_score_codex":0.03208541,"about_ca_topic_score_gemma":0.045848384,"teacher_disagreement_score":0.03208541,"about_ca_system_score_codex":0.003259655,"about_ca_system_score_gemma":0.003971623,"threshold_uncertainty_score":0.063797295},"labels":[],"label_agreement":null},{"id":"W2184821933","doi":"10.1007/s10115-015-0888-6","title":"Mining contentious documents","year":2015,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Information retrieval; Natural language processing; Data science","score_opus":0.037327002593624904,"score_gpt":0.26374361425195963,"score_spread":0.22641661165833474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2184821933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.450472,0.017184662,0.49250874,0.0031568927,0.00064164796,0.0008890263,0.012794656,0.004954062,0.017398316],"genre_scores_gemma":[0.7850077,0.0051479866,0.17455351,0.00034561587,0.0012871067,0.00037442974,0.018356003,0.0005329968,0.0143946195],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99666303,0.00067761954,0.00034346114,0.00083934347,0.0011851962,0.00029134162],"domain_scores_gemma":[0.98930115,0.0057641477,0.000969129,0.0015288823,0.002031063,0.0004055273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001910277,0.0011397614,0.001358361,0.017416166,0.002125046,0.004419041,0.0017564856,0.0016922856,0.0030629917],"category_scores_gemma":[0.015475426,0.00086028816,0.0016420175,0.010902338,0.0009528828,0.0064950846,0.001591307,0.0015510404,0.0024519921],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014691912,0.001016339,0.043413162,0.001971798,0.0010682911,0.0020778982,0.0024652013,0.013914859,0.04632259,0.028107166,0.034559038,0.8236145],"study_design_scores_gemma":[0.00029816298,0.0010269382,0.061127782,0.00081247994,0.0025476175,0.0063098865,0.004351198,0.48786873,0.11366956,0.17431714,0.14740196,0.0002685072],"about_ca_topic_score_codex":0.0022415943,"about_ca_topic_score_gemma":0.0038996928,"teacher_disagreement_score":0.017416166,"about_ca_system_score_codex":0.0010456947,"about_ca_system_score_gemma":0.0018892887,"threshold_uncertainty_score":0.010246754},"labels":[],"label_agreement":null},{"id":"W2185054416","doi":"","title":"The Italica System at TAC 2008 Opinion Summarization Task","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Task (project management); Pyramid (geometry); Sentence; Process (computing); Information retrieval; Relation (database); Natural language processing; Artificial intelligence; Data mining; Programming language; Mathematics; Engineering","score_opus":0.011008175681349104,"score_gpt":0.2252829016983031,"score_spread":0.214274726016954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2185054416","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07506902,0.0035761043,0.41280258,0.004847965,0.0029248279,0.002498002,0.0780756,0.2892676,0.13093837],"genre_scores_gemma":[0.25059915,0.00089293905,0.45034888,0.0015911286,0.00094744394,0.0014352815,0.17527077,0.004976782,0.11393755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99922335,0.00025787574,0.00007978484,0.00024228857,0.00014608701,0.000050630922],"domain_scores_gemma":[0.99847406,0.00047802972,0.0000969856,0.0002885603,0.0005694238,0.000092868155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015163234,0.0012567826,0.0006836307,0.0017010916,0.00061605993,0.0013433472,0.001012317,0.0013588462,0.03630783],"category_scores_gemma":[0.005150238,0.00027193673,0.00056065066,0.00083140266,0.00013953599,0.0019257573,0.0009193784,0.00084382953,0.02967833],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055582874,0.0002074954,0.0011917494,0.0008629225,0.00016805461,0.0003534736,0.00062958937,0.0017222203,0.03731713,0.0031190172,0.45374557,0.50012696],"study_design_scores_gemma":[0.0008902313,0.0008930953,0.008656403,0.00015469354,0.0004617512,0.0010870895,0.0013814837,0.13782881,0.061916202,0.008426901,0.77810556,0.00019780736],"about_ca_topic_score_codex":0.002711569,"about_ca_topic_score_gemma":0.0060829483,"teacher_disagreement_score":0.03630783,"about_ca_system_score_codex":0.000511187,"about_ca_system_score_gemma":0.0006021325,"threshold_uncertainty_score":0.12146181},"labels":[],"label_agreement":null},{"id":"W2185337446","doi":"10.1007/978-3-642-54943-4_13","title":"Looking at Vector Space and Language Models for IR Using Density Matrices","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Space (punctuation); Density matrix; Vector space; Theoretical computer science; Vector space model; Work (physics); Matrix (chemical analysis); Joint (building); Quantum; Algorithm; Algebra over a field; Artificial intelligence; Mathematics; Pure mathematics; Physics","score_opus":0.02338043740809482,"score_gpt":0.25326273663821425,"score_spread":0.22988229923011944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2185337446","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066564796,0.0036027816,0.9795694,0.0033074021,0.00022583171,0.000026370704,0.00025111454,0.00063975883,0.0057209954],"genre_scores_gemma":[0.42949176,0.009622948,0.5095029,0.0025045301,0.0020955927,0.0003379209,0.0018655803,0.0012465945,0.043332256],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972994,0.0016715345,0.00011652928,0.0003697031,0.0003970315,0.00014584896],"domain_scores_gemma":[0.9900034,0.007769228,0.00032825934,0.001038399,0.00065168925,0.0002089895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004313299,0.0011484344,0.0016403157,0.001283863,0.000982444,0.0040984927,0.0022508095,0.002214148,0.010448498],"category_scores_gemma":[0.018179152,0.0011417603,0.002014365,0.0022527967,0.0019842177,0.012437975,0.0018012525,0.0053576827,0.003814927],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000097137985,0.000065817425,0.0007354171,0.00032072476,0.00011183179,0.00008337942,0.00044591862,0.037005648,0.0014828607,0.84967536,0.011300764,0.09867508],"study_design_scores_gemma":[0.000009550823,0.000033523454,0.0001889155,0.000051231607,0.000019935753,0.00008151745,0.00007385235,0.1361949,0.0004715198,0.8580778,0.004765396,0.00003182727],"about_ca_topic_score_codex":0.005057135,"about_ca_topic_score_gemma":0.003727394,"teacher_disagreement_score":0.010448498,"about_ca_system_score_codex":0.0015152019,"about_ca_system_score_gemma":0.0010234022,"threshold_uncertainty_score":0.034953654},"labels":[],"label_agreement":null},{"id":"W2185451889","doi":"","title":"JRC's Participation at TAC 2011: Guided and Multilingual Summarization Tasks","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Readability; Paraphrase; Natural language processing; Task (project management); Redundancy (engineering); Sentence; Grammaticality; Artificial intelligence; Linguistics; Grammar; Programming language","score_opus":0.03353009071716854,"score_gpt":0.282106949434845,"score_spread":0.24857685871767646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2185451889","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6744702,0.0024178685,0.11870444,0.013120101,0.0048507852,0.0063310894,0.064749934,0.064541586,0.05081403],"genre_scores_gemma":[0.5594519,0.0005615616,0.20682894,0.0021629038,0.0014647854,0.003160262,0.1601747,0.008312674,0.05788224],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9914894,0.0035370192,0.0005615886,0.0015697005,0.0017458583,0.001096549],"domain_scores_gemma":[0.97802234,0.0068572364,0.00038004693,0.0032687173,0.008262851,0.0032088258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01179869,0.0024581451,0.0027744872,0.0016324905,0.0030013924,0.0032121409,0.0029324049,0.003663296,0.01218015],"category_scores_gemma":[0.019263411,0.00071409676,0.0015466978,0.0016881757,0.00064451614,0.0030338345,0.0024215449,0.0036372584,0.008770813],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003431699,0.0040648854,0.0031157422,0.001175903,0.00049368164,0.0019755166,0.0065770214,0.010852708,0.06302718,0.0016948173,0.5592549,0.34433594],"study_design_scores_gemma":[0.004001569,0.0075115054,0.0357714,0.00020009534,0.00056430174,0.0023827753,0.0075522787,0.15198532,0.13557065,0.0069647287,0.64625806,0.0012372538],"about_ca_topic_score_codex":0.0136812655,"about_ca_topic_score_gemma":0.020947877,"teacher_disagreement_score":0.0136812655,"about_ca_system_score_codex":0.002084114,"about_ca_system_score_gemma":0.0028563996,"threshold_uncertainty_score":0.062398195},"labels":[],"label_agreement":null},{"id":"W2186748791","doi":"","title":"FRDC's Cross-lingual Entity Linking System at TAC 2013","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Ranking (information retrieval); Entity linking; Acronym; Information retrieval; Lexicon; Population; Artificial intelligence; Cluster analysis; Natural language processing; Heuristic; Knowledge base; Data mining","score_opus":0.009312109205930782,"score_gpt":0.24869923442904746,"score_spread":0.23938712522311667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2186748791","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054570448,0.0036947257,0.28692675,0.0020763536,0.001458882,0.0012908858,0.104283415,0.47071943,0.07497917],"genre_scores_gemma":[0.11785429,0.00075109315,0.4707711,0.0011856584,0.00024621922,0.0008009207,0.37078768,0.011584976,0.026018053],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99448484,0.0011986975,0.00048286436,0.0015715716,0.0017641872,0.0004978435],"domain_scores_gemma":[0.989703,0.0015943986,0.00033429114,0.0033896775,0.0044323225,0.00054631225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006727535,0.0017735569,0.0018257753,0.006889612,0.0027801765,0.0034400846,0.0039023778,0.002734369,0.021757841],"category_scores_gemma":[0.011065855,0.0009416267,0.0012425387,0.0048529822,0.0004812452,0.008731749,0.004381802,0.0029121188,0.03416923],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005969636,0.00062851823,0.0033988196,0.0005812705,0.00023616076,0.0006075661,0.00063262426,0.0033306438,0.013255974,0.0050198548,0.59129626,0.38041538],"study_design_scores_gemma":[0.00045318733,0.00036078403,0.00928699,0.00020616419,0.000359566,0.0018337325,0.0009503616,0.12867841,0.062117387,0.008766693,0.78650945,0.00047722654],"about_ca_topic_score_codex":0.028616136,"about_ca_topic_score_gemma":0.026446255,"teacher_disagreement_score":0.028616136,"about_ca_system_score_codex":0.0021372866,"about_ca_system_score_gemma":0.003636785,"threshold_uncertainty_score":0.072787225},"labels":[],"label_agreement":null},{"id":"W2186752618","doi":"10.1007/978-3-642-30353-1_36","title":"Dsharp: Fast d-DNNF Compilation with sharpSAT","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":108,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Compiler; Exploit; Programming language; Negation; Representation (politics); Theoretical computer science; De facto; Algorithm","score_opus":0.02359578920508502,"score_gpt":0.23590839074299244,"score_spread":0.21231260153790743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2186752618","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018073602,0.0011651171,0.7483744,0.00081114296,0.0009939322,0.00030944057,0.009943276,0.1683753,0.051953815],"genre_scores_gemma":[0.21940929,0.00066046644,0.7114186,0.0010480786,0.000247584,0.0005792118,0.018252695,0.024615064,0.023769068],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9991684,0.0001299106,0.00009021999,0.00028255698,0.00020593582,0.00012300402],"domain_scores_gemma":[0.9983462,0.0006990333,0.00005010108,0.000605468,0.00025005336,0.000049104638],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005354053,0.0018484639,0.0008159629,0.0012359669,0.00090261776,0.0018921006,0.0022985416,0.0009139312,0.0415824],"category_scores_gemma":[0.0029288186,0.0010767345,0.0016211639,0.0017038089,0.00084225903,0.0040747058,0.0028925242,0.0021654025,0.0145127075],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083700113,0.00022944345,0.0012367204,0.0018045594,0.00013688808,0.00027361367,0.00033968902,0.027515614,0.020900898,0.07307726,0.18270952,0.69093883],"study_design_scores_gemma":[0.0007019727,0.00021946702,0.00079303514,0.00042869154,0.00025760377,0.00042128345,0.00038065176,0.285162,0.08309728,0.36800218,0.26035312,0.0001826996],"about_ca_topic_score_codex":0.0042543467,"about_ca_topic_score_gemma":0.010745615,"teacher_disagreement_score":0.0415824,"about_ca_system_score_codex":0.0014250415,"about_ca_system_score_gemma":0.0017156197,"threshold_uncertainty_score":0.13910699},"labels":[],"label_agreement":null},{"id":"W2187008919","doi":"","title":"CUNY-BLENDER TAC-KBP2010 Entity Linking and Slot Filling System Description","year":2010,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Task (project management); Surprise; Entity linking; Sentence; Baseline (sea); Focus (optics); Space (punctuation); Natural language processing; Artificial intelligence; Information retrieval; Knowledge base; Engineering","score_opus":0.01228903924965486,"score_gpt":0.22720201247189317,"score_spread":0.21491297322223832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2187008919","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042011425,0.001318836,0.3273996,0.0023113468,0.0007233361,0.004527711,0.06483368,0.50874645,0.048127674],"genre_scores_gemma":[0.177286,0.00057524274,0.5084798,0.0012590737,0.00024015638,0.0054247426,0.24554399,0.014662592,0.046528477],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99787664,0.0003415363,0.00026035943,0.00067118864,0.00067504763,0.00017516015],"domain_scores_gemma":[0.9963696,0.00074778,0.00015750976,0.0008615539,0.0016100935,0.00025340897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002525468,0.0013244969,0.0013905255,0.002247117,0.0011583787,0.002836688,0.0032984035,0.0013829811,0.0750053],"category_scores_gemma":[0.008975808,0.0011511956,0.00065028085,0.0018604924,0.00039360768,0.004505359,0.002273108,0.0023788873,0.025051815],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084963394,0.00032341882,0.0024596185,0.0013084791,0.00012447473,0.00060003006,0.000738562,0.0064692684,0.02745519,0.006323975,0.58980685,0.3635405],"study_design_scores_gemma":[0.0008154989,0.0007008158,0.0077613466,0.0001895234,0.00018791908,0.00168953,0.0006357299,0.2507798,0.06517176,0.0063730367,0.6653,0.000395051],"about_ca_topic_score_codex":0.022540072,"about_ca_topic_score_gemma":0.020648453,"teacher_disagreement_score":0.0750053,"about_ca_system_score_codex":0.0018349978,"about_ca_system_score_gemma":0.0027952204,"threshold_uncertainty_score":0.25091767},"labels":[],"label_agreement":null},{"id":"W2187175884","doi":"","title":"WHUSUM Participation at TAC 2011 Guided Summarization Track","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Novelty; Computer science; Information retrieval; Multi-document summarization; Task (project management); Track (disk drive); Ranking (information retrieval); Graph; Natural language processing; Theoretical computer science; Engineering","score_opus":0.035802625389545574,"score_gpt":0.26747642530146354,"score_spread":0.23167379991191797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2187175884","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.420622,0.002300414,0.224317,0.010281291,0.0076736687,0.010306589,0.09157043,0.11461145,0.11831717],"genre_scores_gemma":[0.41121817,0.00045571488,0.220023,0.0021915168,0.0018029985,0.0074373335,0.20120703,0.009276718,0.14638759],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9931497,0.0029579203,0.00029685305,0.001053074,0.0017832972,0.00075911876],"domain_scores_gemma":[0.9876072,0.0037340296,0.00034706772,0.0017270269,0.005192252,0.0013924249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009734739,0.0015023213,0.0013115853,0.0020584604,0.0032714335,0.0026758462,0.0024470123,0.002149839,0.021816922],"category_scores_gemma":[0.0165434,0.0005574025,0.0007312932,0.0014677403,0.00053290953,0.0029848248,0.003333932,0.00223988,0.011322707],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003362201,0.0017542269,0.0048075714,0.000925134,0.00029894593,0.0010618382,0.011296022,0.003271241,0.053164,0.0028185262,0.6344247,0.2828156],"study_design_scores_gemma":[0.0014476206,0.0027291696,0.01729912,0.00011434787,0.00029781178,0.0004601652,0.005677734,0.041735534,0.06281417,0.0034958264,0.8635747,0.0003538819],"about_ca_topic_score_codex":0.013593278,"about_ca_topic_score_gemma":0.0268803,"teacher_disagreement_score":0.021816922,"about_ca_system_score_codex":0.0014930584,"about_ca_system_score_gemma":0.0026803813,"threshold_uncertainty_score":0.072984815},"labels":[],"label_agreement":null},{"id":"W2187794223","doi":"","title":"CUNY-UIUC-SRI TAC-KBP2011 Entity Linking System Description","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Focus (optics); NIST; Similarity (geometry); Natural language processing; Track (disk drive); Artificial intelligence; Population; Translation (biology); Simple (philosophy); Entity linking; Information retrieval; Knowledge base; Image (mathematics); Operating system; Medicine; Chemistry","score_opus":0.026483587678195846,"score_gpt":0.2250944414930602,"score_spread":0.19861085381486437,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2187794223","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014850654,0.0032402638,0.3150945,0.0037167084,0.0013105503,0.0036959844,0.14370078,0.4187415,0.09564913],"genre_scores_gemma":[0.037811752,0.001172521,0.33454365,0.0012908394,0.00020037816,0.00287125,0.5656327,0.011616447,0.04486045],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967809,0.0006554867,0.00037770852,0.00069588725,0.0011316111,0.00035835046],"domain_scores_gemma":[0.99468946,0.0006006799,0.00023755932,0.0017247212,0.0023990218,0.00034862827],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037887553,0.0019513115,0.0014869091,0.0054807817,0.0021340745,0.004417271,0.0039530764,0.0025390657,0.03135991],"category_scores_gemma":[0.009142079,0.001242805,0.0011698205,0.005075185,0.00048650414,0.006249409,0.0032207784,0.0032233638,0.050547045],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025643705,0.00030975486,0.0015006261,0.00069061393,0.000078740835,0.00029564992,0.00029854517,0.0038271516,0.008855502,0.007590262,0.74147105,0.2348257],"study_design_scores_gemma":[0.0001564003,0.00019230504,0.0034778833,0.00014562049,0.00011431081,0.0007568909,0.00018591568,0.055758197,0.027011855,0.00395198,0.90805316,0.00019544875],"about_ca_topic_score_codex":0.047905106,"about_ca_topic_score_gemma":0.041979272,"teacher_disagreement_score":0.047905106,"about_ca_system_score_codex":0.0033684194,"about_ca_system_score_gemma":0.0055433125,"threshold_uncertainty_score":0.10490936},"labels":[],"label_agreement":null},{"id":"W2188437695","doi":"","title":"LCC Approaches to Knowledge Base Population at TAC 2010.","year":2010,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Surprise; Knowledge base; Computer science; Task (project management); Context (archaeology); Population; Base (topology); Relation (database); Track (disk drive); Artificial intelligence; Data mining; Engineering; Mathematics; Geography; Systems engineering","score_opus":0.040395806356588994,"score_gpt":0.25342152675725,"score_spread":0.213025720400661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2188437695","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035314737,0.0017382456,0.97990763,0.0012528868,0.00017265188,0.0005360795,0.0011980428,0.006538463,0.005124579],"genre_scores_gemma":[0.03765827,0.0006818833,0.9490033,0.00048451594,0.00020333922,0.0007397172,0.0042290874,0.00085410936,0.006145785],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9881378,0.004615822,0.00069964177,0.0021590774,0.0038907153,0.0004971126],"domain_scores_gemma":[0.9779694,0.008826205,0.00064344506,0.00535725,0.006448758,0.00075492985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014462423,0.0014805588,0.0012184915,0.007769493,0.00245723,0.007502407,0.005311349,0.0027914909,0.015298976],"category_scores_gemma":[0.03788986,0.001604678,0.0017840067,0.0072096754,0.0016680149,0.008601691,0.0053616683,0.0052844784,0.0049559437],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001946192,0.00031939973,0.0015178006,0.00055529695,0.00018126359,0.00024101259,0.0014825454,0.026577925,0.0037333725,0.084507786,0.05421721,0.8264718],"study_design_scores_gemma":[0.00013705075,0.00013628646,0.0013975403,0.00026325652,0.00014380927,0.0004293999,0.0007025736,0.5831633,0.014057069,0.14957067,0.24985847,0.00014053195],"about_ca_topic_score_codex":0.02556498,"about_ca_topic_score_gemma":0.033719778,"teacher_disagreement_score":0.02556498,"about_ca_system_score_codex":0.006834039,"about_ca_system_score_gemma":0.0055077616,"threshold_uncertainty_score":0.076485515},"labels":[],"label_agreement":null},{"id":"W2189491602","doi":"","title":"Lorify: A Knowledge Base from Scratch","year":2012,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Knowledge base; Scalability; Entity linking; Coreference; Cluster analysis; Base (topology); Component (thermodynamics); Task (project management); Population; Scaling; Markov chain; Scratch; Data mining; Artificial intelligence; Theoretical computer science; Machine learning; Database; Programming language; Resolution (logic); Engineering; Mathematics","score_opus":0.01683638553064825,"score_gpt":0.2596367593061633,"score_spread":0.24280037377551505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2189491602","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048690196,0.000595451,0.91115046,0.00069906435,0.00012997766,0.0005520783,0.006466614,0.065236144,0.010301149],"genre_scores_gemma":[0.052811,0.0009647858,0.8982256,0.0010492544,0.00010988141,0.0007096223,0.031399243,0.005046335,0.00968422],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99807334,0.00041682986,0.00014399443,0.00057443493,0.00069045724,0.00010097899],"domain_scores_gemma":[0.9910336,0.003927965,0.00033057676,0.0029443267,0.0013988232,0.00036477155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038916052,0.0012117218,0.0013674494,0.0050092884,0.0014560845,0.0050806613,0.0045611165,0.0017401328,0.016227335],"category_scores_gemma":[0.020502435,0.0012510135,0.0011509955,0.003295984,0.00076700625,0.008725297,0.005396302,0.0031032264,0.014802123],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005036519,0.0005555679,0.0026352266,0.001178117,0.00040312868,0.0007779793,0.00087657414,0.020774491,0.010808639,0.02399621,0.14210734,0.7953831],"study_design_scores_gemma":[0.00024923676,0.00039545822,0.0024902294,0.0007073631,0.00041348126,0.0012578618,0.0008530744,0.32941946,0.028884951,0.0927654,0.5422645,0.00029888866],"about_ca_topic_score_codex":0.006503,"about_ca_topic_score_gemma":0.010103999,"teacher_disagreement_score":0.016227335,"about_ca_system_score_codex":0.00086593523,"about_ca_system_score_gemma":0.0024256203,"threshold_uncertainty_score":0.054285824},"labels":[],"label_agreement":null},{"id":"W2189749715","doi":"","title":"Embedding inference for structured multilabel prediction","year":2015,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Inference; Computer science; Structured prediction; Scalability; Embedding; Machine learning; Margin (machine learning); Representation (politics); Artificial intelligence; Bottleneck; Key (lock); Theoretical computer science","score_opus":0.056424358961094934,"score_gpt":0.3103287019881856,"score_spread":0.25390434302709064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2189749715","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010106214,0.00013937008,0.987566,0.00021741631,0.00002802218,0.000028106111,0.00018289215,0.0010257759,0.0007060915],"genre_scores_gemma":[0.670287,0.00022886203,0.32250026,0.00040156505,0.00017760148,0.00020995976,0.0023520242,0.00047246937,0.0033702487],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99841,0.00065888383,0.00006700948,0.0004849038,0.00026366577,0.00011560568],"domain_scores_gemma":[0.9936651,0.004162259,0.00036288355,0.0011060354,0.0005468098,0.00015707802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021486098,0.0011280758,0.0012315137,0.0008571808,0.0006075025,0.0011794649,0.001936431,0.001283231,0.0042225597],"category_scores_gemma":[0.014020505,0.0006367074,0.0008266644,0.0010369117,0.0010571984,0.003859664,0.0021342074,0.0031046914,0.0013356978],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031374043,0.00022693028,0.002240356,0.0002461713,0.0001295072,0.00017643374,0.0002811796,0.6887457,0.0045439266,0.059489697,0.009608219,0.23399824],"study_design_scores_gemma":[0.000007820779,0.000013504965,0.000090730704,0.000007041961,0.0000057983625,0.000008049914,0.000009199228,0.95084083,0.0006490156,0.04803531,0.00032689763,0.0000058066294],"about_ca_topic_score_codex":0.0024105406,"about_ca_topic_score_gemma":0.005600625,"teacher_disagreement_score":0.0042225597,"about_ca_system_score_codex":0.0009440248,"about_ca_system_score_gemma":0.0009794146,"threshold_uncertainty_score":0.014125884},"labels":[],"label_agreement":null},{"id":"W2192953843","doi":"10.1609/aiide.v9i1.12686","title":"Modeling Autobiographical Memory for Believable Agents","year":2013,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence and Interactive Digital Entertainment","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Recall; Autobiographical memory; Context (archaeology); Memory model; Connectionism; Representation (politics); Human–computer interaction; Cognitive science; Java; Event (particle physics); Artificial intelligence; Cognitive psychology; Psychology; Programming language; Artificial neural network; Shared memory","score_opus":0.06039958024575793,"score_gpt":0.2846761098294408,"score_spread":0.22427652958368288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2192953843","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30721316,0.0001902587,0.68016016,0.00048620783,0.0000593129,0.00007037612,0.00013416763,0.00044495158,0.011241302],"genre_scores_gemma":[0.93959224,0.00016686057,0.05562562,0.00004119896,0.000011490443,0.000099030054,0.00006938147,0.000025492103,0.0043687006],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9999124,0.00002720674,0.0000041426038,0.00002457446,0.000018004126,0.000013657738],"domain_scores_gemma":[0.99964166,0.00018087363,0.000049246442,0.000045815657,0.00005159974,0.00003079754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002497656,0.0003153362,0.00026158307,0.00024618825,0.00032861717,0.00075641874,0.0013934671,0.0008922521,0.0021658621],"category_scores_gemma":[0.0011221087,0.0002711926,0.00039577583,0.00015996577,0.00054208725,0.0014616033,0.0005851767,0.00072074734,0.00023260631],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006093106,0.00005122031,0.0008323847,0.000038356462,0.00002824565,0.00011815219,0.0002037076,0.9597036,0.0032619555,0.029231302,0.00020258588,0.0062675998],"study_design_scores_gemma":[0.0000051561115,0.000011362116,0.00006569148,0.0000016499203,0.0000044677017,0.000009886576,0.000009232368,0.99503183,0.00031048557,0.0042795823,0.00026804564,0.0000025740549],"about_ca_topic_score_codex":0.009181167,"about_ca_topic_score_gemma":0.009385357,"teacher_disagreement_score":0.009181167,"about_ca_system_score_codex":0.0008292051,"about_ca_system_score_gemma":0.0006241814,"threshold_uncertainty_score":0.018255472},"labels":[],"label_agreement":null},{"id":"W2197022017","doi":"10.33011/lilt.v12i.1375","title":"Distinguishing Voices in The Waste Land using Computational Stylistics","year":2015,"lang":"en","type":"article","venue":"Linguistic Issues in Language Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Stylistics; Viewpoints; Computer science; Interpretation (philosophy); Poetry; Cluster analysis; Natural language processing; Linguistics; Citation; Perspective (graphical); Segmentation; Artificial intelligence; Computational linguistics; Art; Visual arts; Philosophy","score_opus":0.03122741613727188,"score_gpt":0.3270873528477606,"score_spread":0.2958599367104887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2197022017","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57600087,0.00047818586,0.40306604,0.0011748739,0.00010224565,0.00011937334,0.00050185446,0.00042505245,0.01813149],"genre_scores_gemma":[0.93891907,0.00014824697,0.05861402,0.00006945943,0.000047820376,0.000044776978,0.00035567654,0.00010508115,0.0016959561],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99878067,0.0006740824,0.00007058517,0.00022461075,0.00018914157,0.000060899947],"domain_scores_gemma":[0.9945529,0.0040096226,0.00048339964,0.0005300888,0.00031812213,0.000105895844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015464395,0.00035234084,0.00031975243,0.0029270728,0.0012417004,0.0042151897,0.00047861013,0.00058524864,0.001393097],"category_scores_gemma":[0.0084510995,0.00024372988,0.0005504133,0.0018162645,0.0022189908,0.0035134626,0.0016160919,0.0009054746,0.00052861567],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074034644,0.00012936861,0.070199385,0.0005870535,0.00015065885,0.0016015542,0.097085424,0.040197674,0.0448684,0.3046901,0.0077868546,0.43196326],"study_design_scores_gemma":[0.000043895776,0.00014250226,0.05349884,0.00030321206,0.00008684561,0.0019184987,0.03538497,0.49443337,0.023912072,0.32058647,0.06949225,0.00019708966],"about_ca_topic_score_codex":0.0013315768,"about_ca_topic_score_gemma":0.002537191,"teacher_disagreement_score":0.0042151897,"about_ca_system_score_codex":0.00060589286,"about_ca_system_score_gemma":0.0004723861,"threshold_uncertainty_score":0.0081784725},"labels":[],"label_agreement":null},{"id":"W2198432679","doi":"","title":"A Belief Revision Approach to Textual Entailment Recognition","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Textual entailment; Logical consequence; Computer science; Natural language processing; Artificial intelligence; Belief revision","score_opus":0.06861730497020792,"score_gpt":0.2537095289288492,"score_spread":0.18509222395864128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2198432679","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008954642,0.00018398595,0.98675585,0.0003713444,0.000049066075,0.00008414643,0.00012390423,0.0014423282,0.0020347773],"genre_scores_gemma":[0.37107402,0.00024913982,0.62264925,0.00027008328,0.00024002217,0.00015573317,0.000554949,0.0001461129,0.0046606013],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977107,0.00083906227,0.00017002715,0.00046994223,0.0006835732,0.00012667802],"domain_scores_gemma":[0.9954515,0.0022181782,0.00029110847,0.0008104134,0.0010643258,0.00016453088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026945279,0.0005948378,0.0007453208,0.0015347828,0.00074106123,0.0023527883,0.0019378915,0.0012506966,0.004469394],"category_scores_gemma":[0.010454849,0.00042293884,0.0014022826,0.00089732744,0.0011755868,0.0032069173,0.0012872037,0.002057634,0.0013019231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056619645,0.00035714585,0.0024414253,0.00056928216,0.00039181806,0.0004220662,0.0018489907,0.03615409,0.049245786,0.12895916,0.010244669,0.7687993],"study_design_scores_gemma":[0.000059407466,0.00024608802,0.0018341963,0.00006726543,0.00016932348,0.0004895567,0.00025554185,0.831448,0.026508726,0.1296203,0.009172915,0.00012871785],"about_ca_topic_score_codex":0.002281382,"about_ca_topic_score_gemma":0.0027079002,"teacher_disagreement_score":0.004469394,"about_ca_system_score_codex":0.0006973941,"about_ca_system_score_gemma":0.0006506001,"threshold_uncertainty_score":0.014951587},"labels":[],"label_agreement":null},{"id":"W2208988332","doi":"10.33011/lilt.v12i.1373","title":"Literature Lifts Up Computational Linguistics","year":2015,"lang":"en","type":"article","venue":"Linguistic Issues in Language Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; National Research Council Canada","funders":"","keywords":"Computational linguistics; Computer science; Linguistics; Cognitive linguistics; Natural language processing; Applied linguistics; Artificial intelligence; Philosophy; Psychology","score_opus":0.01660944893564042,"score_gpt":0.30438343092791736,"score_spread":0.28777398199227694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2208988332","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010540724,0.07906156,0.24537845,0.17463899,0.027950617,0.0001062184,0.0018323716,0.0019700686,0.45852107],"genre_scores_gemma":[0.4978043,0.06771634,0.1347487,0.04004401,0.07371088,0.00059683603,0.0053682653,0.003687917,0.17632276],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9937052,0.0032239566,0.0004007976,0.0010071368,0.0013560883,0.00030687914],"domain_scores_gemma":[0.98168266,0.011190315,0.0006195306,0.003803428,0.0019617563,0.00074239535],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0044186637,0.0009231502,0.0010692206,0.0102779735,0.004700966,0.009781507,0.0014763366,0.0020573167,0.022662535],"category_scores_gemma":[0.02133617,0.00077937567,0.0018804632,0.0068051596,0.008298561,0.017552264,0.0082548885,0.0065647406,0.0098531395],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011814079,0.000010148528,0.00016826723,0.00015662647,0.000017477587,0.00008445506,0.001314331,0.00022478156,0.00011056915,0.91271824,0.052803747,0.032379597],"study_design_scores_gemma":[0.000004065558,0.0000043689847,0.00016839776,0.0002255276,0.0000074087034,0.00011750981,0.0005855967,0.0006366546,0.0000876807,0.5891171,0.40903437,0.000011312438],"about_ca_topic_score_codex":0.0024644323,"about_ca_topic_score_gemma":0.0032391285,"teacher_disagreement_score":0.989722,"about_ca_system_score_codex":0.0035657366,"about_ca_system_score_gemma":0.0027717513,"threshold_uncertainty_score":0.07581377},"labels":[],"label_agreement":null},{"id":"W2220346254","doi":"10.1109/wcitca.2015.7367038","title":"PICO extraction by combining the robustness of machine-learning methods with the rule-based methods","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Robustness (evolution); Computer science; Artificial intelligence; Feature extraction; Machine learning; Data mining; Pattern recognition (psychology)","score_opus":0.061727591370695695,"score_gpt":0.3625301398504605,"score_spread":0.30080254847976484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2220346254","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007152223,0.00089628884,0.9863063,0.00019584944,0.000074053176,0.0002946631,0.00032685327,0.002684919,0.0020688302],"genre_scores_gemma":[0.16361965,0.00092898385,0.82883465,0.00020241934,0.00036967915,0.000502954,0.0027152176,0.0006143328,0.0022120608],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99484396,0.0012160877,0.0008149872,0.0014266327,0.0014158722,0.00028243274],"domain_scores_gemma":[0.9894013,0.006572056,0.0006880096,0.0014219121,0.0017425928,0.00017411988],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0052939816,0.001741263,0.0019053881,0.008830152,0.000900644,0.0035593435,0.0018372237,0.0015455253,0.0022477864],"category_scores_gemma":[0.018828794,0.0007611339,0.0021475235,0.0036618276,0.00094557035,0.0040507563,0.0025725693,0.0016536438,0.0023324047],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016686131,0.000128959,0.004827855,0.0005252475,0.000224322,0.00034804308,0.00042842847,0.02865367,0.0144420555,0.007432308,0.0061218673,0.9367003],"study_design_scores_gemma":[0.00006417454,0.00009240517,0.0066041714,0.00019155418,0.00036134143,0.0003663418,0.0003033316,0.9063676,0.02424962,0.03183136,0.029449059,0.0001191135],"about_ca_topic_score_codex":0.0034293563,"about_ca_topic_score_gemma":0.0036261498,"teacher_disagreement_score":0.99470603,"about_ca_system_score_codex":0.000663735,"about_ca_system_score_gemma":0.0015764852,"threshold_uncertainty_score":0.027997553},"labels":[],"label_agreement":null},{"id":"W2229046451","doi":"10.1007/978-3-319-06605-9_13","title":"Topic Modeling Using Collapsed Typed Dependency Relations","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Perplexity; Computer science; Bigram; Natural language processing; Dependency (UML); Topic model; Artificial intelligence; Thematic structure; Coherence (philosophical gambling strategy); Information retrieval; Language model; Programming language","score_opus":0.036461201158759435,"score_gpt":0.2591578996023134,"score_spread":0.22269669844355394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2229046451","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004410884,0.000619801,0.9899573,0.00020790928,0.00007487725,0.00006209597,0.0007083421,0.001631733,0.0023270287],"genre_scores_gemma":[0.2634887,0.0021188352,0.70930254,0.00018400466,0.00037589198,0.00050704635,0.006322336,0.0019918364,0.0157088],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975298,0.0010947756,0.00020928218,0.0005436535,0.00045964847,0.00016285396],"domain_scores_gemma":[0.993171,0.0048727444,0.00022232364,0.00081564084,0.00073280965,0.00018545586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002687901,0.0010377903,0.00096349913,0.0033293955,0.0012967113,0.0030919325,0.0019376646,0.0010114103,0.008744306],"category_scores_gemma":[0.011973874,0.001331904,0.0029707304,0.0047230404,0.0005377813,0.006909097,0.0021291655,0.0026441629,0.0040189587],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052805257,0.00028227363,0.005431399,0.00085962785,0.0005525859,0.0004042843,0.0022142509,0.10116892,0.008480315,0.2293148,0.038226575,0.6125369],"study_design_scores_gemma":[0.000031956537,0.00003258931,0.0008870652,0.000096433505,0.00019125783,0.00019629052,0.0001857208,0.7932708,0.0042591942,0.17479973,0.025991002,0.00005786747],"about_ca_topic_score_codex":0.008221046,"about_ca_topic_score_gemma":0.012102267,"teacher_disagreement_score":0.008744306,"about_ca_system_score_codex":0.0012773606,"about_ca_system_score_gemma":0.0015644052,"threshold_uncertainty_score":0.029252589},"labels":[],"label_agreement":null},{"id":"W2234607897","doi":"10.1080/17470218.2015.1130068","title":"Applying an exemplar model to an implicit rule-learning task: Implicit learning of semantic structure","year":2016,"lang":"en","type":"article","venue":"Quarterly Journal of Experimental Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Implicit learning; Task (project management); Sequence learning; Artificial intelligence; Natural language processing; Computer science; Cognitive psychology; Recall; Cognition; Psychology; Machine learning","score_opus":0.024279793044816735,"score_gpt":0.33620627414303966,"score_spread":0.3119264810982229,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2234607897","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6716325,0.00009744042,0.32283214,0.0005393893,0.00011447494,0.0001665985,0.00026841759,0.000501744,0.0038471918],"genre_scores_gemma":[0.89008737,0.00012245319,0.10782353,0.00011151298,0.000023650775,0.0001862518,0.0003728501,0.000056714274,0.0012157061],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99944824,0.00021807916,0.00004584114,0.00016460678,0.00009498478,0.00002827295],"domain_scores_gemma":[0.99407417,0.004151456,0.00032578083,0.0011213215,0.00020374134,0.00012362629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002009462,0.0005124713,0.00077278377,0.00032675554,0.0003643123,0.0013722203,0.0021523396,0.0015128115,0.0030599497],"category_scores_gemma":[0.012768746,0.00042769054,0.0008316796,0.00045843044,0.0009183808,0.0038659812,0.0014034449,0.0021542194,0.000534244],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038882415,0.0039277375,0.033937477,0.001097745,0.00076947984,0.0009772717,0.003515168,0.4227213,0.08895781,0.0925414,0.005272478,0.34239393],"study_design_scores_gemma":[0.00012784664,0.00043242495,0.0015034544,0.00001335268,0.000052946074,0.00016504114,0.00006936194,0.9470904,0.010566906,0.03900591,0.0009376511,0.000034622426],"about_ca_topic_score_codex":0.0012374081,"about_ca_topic_score_gemma":0.0012467603,"teacher_disagreement_score":0.0030599497,"about_ca_system_score_codex":0.0004891458,"about_ca_system_score_gemma":0.0004445488,"threshold_uncertainty_score":0.01062721},"labels":[],"label_agreement":null},{"id":"W2235613556","doi":"10.1007/978-3-642-32790-2_25","title":"Supervised Distributional Semantic Relatedness","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Word (group theory); Weighting; Semantic similarity; Natural language processing; Similarity (geometry); Artificial intelligence; Context (archaeology); Measure (data warehouse); Construct (python library); Similarity measure; Word Association; Variety (cybernetics); Association (psychology); Data mining; Mathematics; Image (mathematics)","score_opus":0.022791926180044872,"score_gpt":0.2386288397965285,"score_spread":0.2158369136164836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2235613556","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02145773,0.002558544,0.93565893,0.0008265594,0.0004905611,0.00012594963,0.0017331468,0.0042965496,0.03285206],"genre_scores_gemma":[0.41894865,0.0027854487,0.49737588,0.0005249872,0.0010454915,0.00038992052,0.02346977,0.0018362602,0.053623587],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979907,0.00057599164,0.000121068195,0.0007067428,0.0005006789,0.000104815175],"domain_scores_gemma":[0.9983473,0.00048165646,0.00009496886,0.00066754455,0.00033793718,0.000070555885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011844644,0.00089779816,0.0011014258,0.002353404,0.0011038639,0.0017498861,0.0015706358,0.0012568125,0.008916818],"category_scores_gemma":[0.0048766737,0.0005275023,0.0011596306,0.0028654796,0.0007646453,0.0045980453,0.0028512583,0.002032599,0.008621647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017126734,0.00027674765,0.0009488498,0.00029037552,0.00010527372,0.00009259105,0.00020651337,0.009861425,0.014700522,0.09187036,0.05590547,0.8255705],"study_design_scores_gemma":[0.000047910195,0.00016632487,0.0033059502,0.00014378579,0.00013174406,0.0010953385,0.00028735047,0.33722776,0.019234885,0.53490174,0.10337292,0.00008424964],"about_ca_topic_score_codex":0.0005922867,"about_ca_topic_score_gemma":0.0016190953,"teacher_disagreement_score":0.008916818,"about_ca_system_score_codex":0.00060854846,"about_ca_system_score_gemma":0.0010071563,"threshold_uncertainty_score":0.02982974},"labels":[],"label_agreement":null},{"id":"W2247119764","doi":"10.1609/aaai.v25i1.7917","title":"Learning Structured Embeddings of Knowledge Bases","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":861,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"WordNet; Computer science; Artificial intelligence; Process (computing); Natural language processing; Information retrieval; Knowledge extraction; Word (group theory); Annotation; Programming language","score_opus":0.036990518651056306,"score_gpt":0.2521480397486367,"score_spread":0.2151575210975804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2247119764","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07634914,0.0005315537,0.9178073,0.0006459523,0.00008850132,0.00008713041,0.0009658724,0.001539409,0.0019850705],"genre_scores_gemma":[0.6783291,0.00080368435,0.31045157,0.00025918064,0.00011750939,0.0002863142,0.005705477,0.00019888011,0.0038482768],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904555,0.0003387778,0.000083359686,0.0002857021,0.00017129778,0.00007533197],"domain_scores_gemma":[0.99563783,0.0026932422,0.00030660574,0.00067133026,0.0005792411,0.0001118383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011804411,0.0008318274,0.0006526685,0.001773969,0.00035045296,0.0015893066,0.0013076476,0.0011849982,0.0018471808],"category_scores_gemma":[0.01101155,0.0006022921,0.0007576907,0.0015869355,0.000750072,0.0050769025,0.0018542195,0.0016949495,0.00067115534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003316166,0.0003002608,0.0054345145,0.00048412057,0.00019077485,0.00025821253,0.00079692685,0.29023632,0.0060454714,0.06519171,0.009112429,0.6216176],"study_design_scores_gemma":[0.000016853366,0.00006237337,0.0003591555,0.00004145847,0.000023840992,0.000031068135,0.0001043099,0.9319354,0.0016238699,0.06378744,0.0020006255,0.0000136221215],"about_ca_topic_score_codex":0.0020664288,"about_ca_topic_score_gemma":0.0036361974,"teacher_disagreement_score":0.0020664288,"about_ca_system_score_codex":0.0007710074,"about_ca_system_score_gemma":0.00070035056,"threshold_uncertainty_score":0.0062428117},"labels":[],"label_agreement":null},{"id":"W2248901308","doi":"10.1109/ictai.2015.41","title":"Lexical Semantic Relatedness for Twitter Analytics","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Toronto Metropolitan University","funders":"","keywords":"WordNet; Microblogging; Computer science; Semantic similarity; Social media; Latent semantic analysis; Semantics (computer science); Information retrieval; Natural language processing; Construct (python library); Graph; Artificial intelligence; World Wide Web","score_opus":0.12566804880840274,"score_gpt":0.31303714909533387,"score_spread":0.18736910028693113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2248901308","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051285967,0.0014626462,0.9283378,0.0013691962,0.00009962953,0.00038251866,0.0030207974,0.0041571297,0.009884379],"genre_scores_gemma":[0.62285984,0.0008398733,0.36676708,0.00025676136,0.00028813418,0.0005911693,0.005822447,0.00041158817,0.0021629683],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963904,0.00186151,0.00027023852,0.00062024745,0.00073318055,0.00012439363],"domain_scores_gemma":[0.9932226,0.0040130657,0.00084911764,0.0009349563,0.0007552165,0.00022508092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032686095,0.0009843422,0.00090437866,0.0076600434,0.0012087091,0.002302235,0.00096577615,0.0012089993,0.005232102],"category_scores_gemma":[0.024113422,0.00039617674,0.0012032286,0.006957787,0.0007406048,0.0071791857,0.0022874267,0.0012452987,0.002538491],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045838437,0.0005763471,0.019168315,0.00093453174,0.0003230946,0.00045599331,0.0021956272,0.066234894,0.010699113,0.24405041,0.01933022,0.63557315],"study_design_scores_gemma":[0.000030704243,0.00012101958,0.0078014596,0.00012694843,0.00007773291,0.00025741995,0.000971634,0.6252154,0.002694975,0.344608,0.018021481,0.00007314235],"about_ca_topic_score_codex":0.0025830346,"about_ca_topic_score_gemma":0.0026893022,"teacher_disagreement_score":0.0076600434,"about_ca_system_score_codex":0.001077364,"about_ca_system_score_gemma":0.0010270644,"threshold_uncertainty_score":0.017503083},"labels":[],"label_agreement":null},{"id":"W2250571015","doi":"","title":"Towards Automatic Topical Question Generation","year":2012,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science","score_opus":0.09053207617422618,"score_gpt":0.3518676068766896,"score_spread":0.2613355307024634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250571015","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023131201,0.00077882735,0.9401515,0.001647788,0.00073752203,0.0004541569,0.002308735,0.023192486,0.007597761],"genre_scores_gemma":[0.21969523,0.00046158046,0.7528078,0.000531251,0.0005123971,0.00052931416,0.013795103,0.0019235619,0.009743728],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967039,0.0014957131,0.0002412061,0.00076310063,0.0005378001,0.0002583547],"domain_scores_gemma":[0.9938467,0.0029065094,0.00016948419,0.0008881327,0.0019609334,0.00022824388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034923558,0.0012508513,0.001346172,0.003915186,0.0015461795,0.0035698882,0.0017114736,0.0022302617,0.017081946],"category_scores_gemma":[0.010108998,0.0009436179,0.0016439509,0.001714645,0.0007190648,0.0047794003,0.003882206,0.0032538164,0.012831376],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004678308,0.000340961,0.0020822452,0.000614998,0.000114529605,0.0003783944,0.0010807791,0.0067855907,0.0718988,0.036154035,0.080779396,0.7993025],"study_design_scores_gemma":[0.00014868853,0.00017336647,0.0014616614,0.000121853096,0.00017730739,0.00053162826,0.000894799,0.74154335,0.07270249,0.09500056,0.08717552,0.000068821064],"about_ca_topic_score_codex":0.0014722939,"about_ca_topic_score_gemma":0.0021231268,"teacher_disagreement_score":0.017081946,"about_ca_system_score_codex":0.00093499693,"about_ca_system_score_gemma":0.001736217,"threshold_uncertainty_score":0.05714482},"labels":[],"label_agreement":null},{"id":"W2250669704","doi":"","title":"Probabilistic Domain Modelling With Contextualized Distributional Semantic Vectors","year":2013,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Automatic summarization; Computer science; Distributional semantics; Artificial intelligence; Probabilistic logic; Natural language processing; Domain (mathematical analysis); Semantics (computer science); Generative grammar; Generative model; Hidden Markov model; Semantic space; Machine learning; Semantic similarity; Mathematics","score_opus":0.01520811403659151,"score_gpt":0.23146867420578987,"score_spread":0.21626056016919837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250669704","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007571125,0.00012845221,0.9898398,0.00017063208,0.000029020108,0.000044824923,0.00031459288,0.0006781381,0.0012234147],"genre_scores_gemma":[0.444993,0.0005714953,0.5437826,0.00021749867,0.0001460398,0.00042374505,0.0037635474,0.0004956953,0.0056064287],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986442,0.000603499,0.000084420564,0.000385546,0.00019933764,0.00008312689],"domain_scores_gemma":[0.9969458,0.001975573,0.00020771474,0.00047115627,0.00032075722,0.00007901756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017737547,0.00070363673,0.0006232508,0.0015632474,0.000508959,0.0016041809,0.0018844877,0.0010479225,0.0032828178],"category_scores_gemma":[0.006248565,0.00058506994,0.0014262212,0.0018680362,0.0007194322,0.0033144355,0.0016312297,0.0019098513,0.0015155532],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019273399,0.00015056673,0.0029732094,0.00026465362,0.00014487238,0.00029207394,0.00071454915,0.5781695,0.005690348,0.2160835,0.0074942303,0.18782976],"study_design_scores_gemma":[0.000011477206,0.000018068013,0.00023428767,0.00001482022,0.000015845433,0.000058856476,0.000046221656,0.91741914,0.0012502226,0.077647336,0.0032683017,0.000015398002],"about_ca_topic_score_codex":0.0026657472,"about_ca_topic_score_gemma":0.005944633,"teacher_disagreement_score":0.0032828178,"about_ca_system_score_codex":0.0008596629,"about_ca_system_score_gemma":0.0010589861,"threshold_uncertainty_score":0.010982096},"labels":[],"label_agreement":null},{"id":"W2250749132","doi":"10.3115/v1/w14-4318","title":"Extractive Summarization and Dialogue Act Modeling on Email Threads: An Integrated Probabilistic Approach","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Conversation; Task (project management); Probabilistic logic; Artificial intelligence; Natural language processing; Statistical model; Graphical model; Machine learning; Linguistics","score_opus":0.03787629057484056,"score_gpt":0.2470056146718206,"score_spread":0.20912932409698004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250749132","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005173787,0.00028276662,0.9929007,0.00017611225,0.000024639867,0.000050836694,0.00015448239,0.00090103026,0.00033565966],"genre_scores_gemma":[0.344034,0.0008521006,0.6480673,0.00020418891,0.0005212285,0.00053056556,0.001764655,0.00039602804,0.003629918],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99720097,0.001197684,0.00022636686,0.0007708736,0.00048200492,0.00012206313],"domain_scores_gemma":[0.9937423,0.004091358,0.0007261631,0.0005371896,0.00074686337,0.00015611945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003062185,0.0013271195,0.0013191572,0.002541464,0.00064838777,0.001857014,0.0020650544,0.0014810249,0.0016119021],"category_scores_gemma":[0.0090845525,0.00086692674,0.0017568626,0.001486265,0.00063918746,0.003425444,0.0015014904,0.0017366814,0.00095732766],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000538356,0.00033785836,0.0053736684,0.00077368756,0.00052994164,0.00031619854,0.0015113255,0.3428813,0.017841687,0.031957787,0.006731752,0.5912065],"study_design_scores_gemma":[0.000013023118,0.00005825368,0.00074698596,0.00002077933,0.00006489193,0.000052314263,0.000057637524,0.97821724,0.0018477925,0.016529901,0.0023683796,0.000022666653],"about_ca_topic_score_codex":0.0031147404,"about_ca_topic_score_gemma":0.005438184,"teacher_disagreement_score":0.0031147404,"about_ca_system_score_codex":0.00089997577,"about_ca_system_score_gemma":0.001304152,"threshold_uncertainty_score":0.016194582},"labels":[],"label_agreement":null},{"id":"W2250777749","doi":"","title":"A High-Precision Approach to Detecting Hedges and their Scopes","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Natural language processing; Scope (computer science); Artificial intelligence; Vagueness; Categorization; Software portability; Weighting; Sentence; Parsing; Task (project management); Domain (mathematical analysis); Extensibility; Programming language; Fuzzy logic","score_opus":0.021178327233485134,"score_gpt":0.23399842457739647,"score_spread":0.21282009734391133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250777749","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024517382,0.00044962,0.9670026,0.00029916211,0.00003711864,0.00019069128,0.0006314873,0.005279405,0.001592525],"genre_scores_gemma":[0.3117805,0.00021571707,0.68389124,0.00011638939,0.00010483916,0.00016836313,0.0014789096,0.00028220573,0.001961877],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964516,0.0008054121,0.00036498488,0.0010824442,0.0010628297,0.00023270892],"domain_scores_gemma":[0.99122405,0.004031001,0.0007500543,0.0019842428,0.001795872,0.0002148413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005091635,0.0015182832,0.001239113,0.006170527,0.0014138084,0.0038003486,0.0025355658,0.0022035607,0.0032401432],"category_scores_gemma":[0.013144026,0.0010263118,0.0013959004,0.0025216711,0.0009784204,0.004992075,0.002778802,0.0023174293,0.0018087038],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002669943,0.0002932655,0.012936588,0.0005294434,0.00028581763,0.0003202737,0.0012140005,0.026939606,0.075986,0.015522568,0.008112397,0.85759294],"study_design_scores_gemma":[0.000045565215,0.00016854398,0.007803967,0.00008295586,0.00019406677,0.0006002478,0.0003181112,0.9036649,0.045006294,0.034769565,0.0072496333,0.00009615113],"about_ca_topic_score_codex":0.0033967756,"about_ca_topic_score_gemma":0.005566878,"teacher_disagreement_score":0.006170527,"about_ca_system_score_codex":0.0010287218,"about_ca_system_score_gemma":0.0016481228,"threshold_uncertainty_score":0.026927471},"labels":[],"label_agreement":null},{"id":"W2250818554","doi":"10.18653/v1/d15-1163","title":"Improving Statistical Machine Translation with a Multilingual Paraphrase Database","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Paraphrase; Computer science; Machine translation; Natural language processing; Artificial intelligence; Translation (biology); Vocabulary; Machine translation system","score_opus":0.05416993627736382,"score_gpt":0.29435074691882174,"score_spread":0.2401808106414579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250818554","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08912085,0.001476219,0.88159925,0.0006346902,0.00015701042,0.00027786964,0.002612498,0.019481717,0.004639914],"genre_scores_gemma":[0.3760741,0.0008508107,0.599234,0.00031215933,0.00023057149,0.0003945334,0.016834687,0.0013897803,0.004679303],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966827,0.0016788675,0.00024967862,0.0006051244,0.00066126615,0.00012229568],"domain_scores_gemma":[0.9913805,0.004358184,0.00039420117,0.001637632,0.0020882704,0.00014110822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002845697,0.0014693828,0.0012661043,0.0049002273,0.0008409763,0.0016733143,0.0019720222,0.0012555156,0.0034046748],"category_scores_gemma":[0.01239982,0.00064275594,0.0011048181,0.006291559,0.0004576526,0.003421769,0.0021749341,0.0018678245,0.004038498],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040312854,0.0006174072,0.0032250853,0.00041671295,0.00033544298,0.00034946826,0.0002637795,0.061712988,0.030691028,0.005141747,0.023614356,0.8732289],"study_design_scores_gemma":[0.0000791748,0.00014296522,0.0018023007,0.000024938821,0.000097846925,0.00019882669,0.00013337217,0.9666578,0.01684563,0.00674335,0.0072331717,0.000040570987],"about_ca_topic_score_codex":0.00665332,"about_ca_topic_score_gemma":0.010738475,"teacher_disagreement_score":0.00665332,"about_ca_system_score_codex":0.00071680086,"about_ca_system_score_gemma":0.0015327936,"threshold_uncertainty_score":0.015049636},"labels":[],"label_agreement":null},{"id":"W2250834857","doi":"10.18653/v1/w15-4307","title":"NRC: Infused Phrase Vectors for Named Entity Recognition in Twitter","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Phrase; Computer science; Discriminative model; Natural language processing; Artificial intelligence; Hidden Markov model; Task (project management); Word (group theory); Named-entity recognition; Set (abstract data type); Training set; Speech recognition; Information retrieval; Linguistics; Engineering","score_opus":0.12851933362443255,"score_gpt":0.296213245727793,"score_spread":0.16769391210336046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250834857","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.087128304,0.001893066,0.70563775,0.002167459,0.0021935157,0.0010974536,0.0641446,0.124921806,0.0108160535],"genre_scores_gemma":[0.23448738,0.0008472471,0.61034733,0.0005987239,0.0004986107,0.0011103086,0.12878355,0.0036751386,0.01965163],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977049,0.0007028401,0.00019788739,0.0006617812,0.00055830093,0.00017428429],"domain_scores_gemma":[0.9963529,0.0013567011,0.00017847265,0.0010614765,0.00085758336,0.0001927945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032234602,0.0020397971,0.00097178953,0.001899267,0.0008707494,0.0015288185,0.001931701,0.0018213217,0.015120379],"category_scores_gemma":[0.01114043,0.0005003337,0.0010307393,0.0020454447,0.0004021364,0.0058519244,0.003210402,0.0020574185,0.016635584],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011376921,0.00044679464,0.005197456,0.0007785191,0.000234707,0.0004920005,0.0005117538,0.01105315,0.036366433,0.0053375424,0.24528483,0.6931591],"study_design_scores_gemma":[0.00021724768,0.0006099089,0.007849402,0.00016503801,0.000113204194,0.0008641611,0.0006305598,0.8031401,0.05768151,0.014094981,0.1144239,0.00021001277],"about_ca_topic_score_codex":0.0051712575,"about_ca_topic_score_gemma":0.007497221,"teacher_disagreement_score":0.015120379,"about_ca_system_score_codex":0.000649662,"about_ca_system_score_gemma":0.0009831735,"threshold_uncertainty_score":0.050582707},"labels":[],"label_agreement":null},{"id":"W2250838119","doi":"","title":"Application of the Tightness Continuum Measure to Chinese Information Retrieval","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Segmentation; Artificial intelligence; Text segmentation; Word (group theory); Natural language processing; Measure (data warehouse); Pattern recognition (psychology); Information retrieval; Data mining; Mathematics","score_opus":0.005421990263527392,"score_gpt":0.2200545502786905,"score_spread":0.21463256001516312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250838119","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25444093,0.0041428446,0.7308286,0.0004745726,0.0000775278,0.0003111878,0.00033480336,0.0013951837,0.007994339],"genre_scores_gemma":[0.871751,0.00074533303,0.12584649,0.000086272266,0.00015756476,0.00020026829,0.0004008511,0.00010415824,0.00070804235],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961665,0.0014279,0.00041759692,0.0005727866,0.0011825827,0.00023251191],"domain_scores_gemma":[0.9909992,0.0056605823,0.00091566937,0.0010093788,0.0010411219,0.00037407334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052856817,0.0006700966,0.0011033287,0.008254373,0.0014352826,0.0019160606,0.00093608384,0.0007075297,0.0012111524],"category_scores_gemma":[0.019534389,0.00044477877,0.00073541305,0.0072665988,0.0022797317,0.0036825342,0.0022045623,0.001122516,0.00024742528],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006387658,0.0002755902,0.028707199,0.00083519984,0.00035887668,0.0005708788,0.003242034,0.14671981,0.028253753,0.11152354,0.0048593353,0.674015],"study_design_scores_gemma":[0.00006333735,0.00045187288,0.03286207,0.0000868572,0.00013842147,0.00049901113,0.0006808121,0.8525136,0.010564952,0.0959886,0.0059809512,0.00016941647],"about_ca_topic_score_codex":0.007301911,"about_ca_topic_score_gemma":0.0037821482,"teacher_disagreement_score":0.008254373,"about_ca_system_score_codex":0.0019472158,"about_ca_system_score_gemma":0.0015541258,"threshold_uncertainty_score":0.027953684},"labels":[],"label_agreement":null},{"id":"W2250863190","doi":"10.3115/v1/d14-1051","title":"An I-vector Based Approach to Compact Multi-Granularity Topic Spaces Representation of Textual Documents","year":2014,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Agence Nationale de la Recherche","keywords":"Granularity; Computer science; Representation (politics); Information retrieval; Vector space; Natural language processing; Artificial intelligence; Mathematics; Programming language; Pure mathematics","score_opus":0.06669668175130425,"score_gpt":0.33794628761034246,"score_spread":0.27124960585903823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250863190","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01355212,0.0037542472,0.97493005,0.00047697386,0.0002477222,0.00016758239,0.0021997376,0.0028736445,0.001798062],"genre_scores_gemma":[0.27794358,0.003073081,0.6980144,0.00025485858,0.0006929697,0.00067185843,0.012266494,0.0005499924,0.006532707],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99854976,0.00044007364,0.00015844124,0.00037259166,0.00032275653,0.00015638178],"domain_scores_gemma":[0.99752,0.0010963357,0.00019981252,0.0003176206,0.00071764324,0.00014861824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001367663,0.0009400513,0.0011482708,0.0050214976,0.00070453284,0.0024070723,0.0015890956,0.0011449504,0.004096368],"category_scores_gemma":[0.005077295,0.00038036902,0.001312402,0.006415865,0.00045069604,0.003534524,0.0017016707,0.001744996,0.0027745415],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080521597,0.00028600616,0.0020775064,0.0005780964,0.00017579376,0.00013476776,0.0008257534,0.027442178,0.01486517,0.024251698,0.025045987,0.90351194],"study_design_scores_gemma":[0.00007986888,0.00024473088,0.0022767547,0.0001127056,0.00012727398,0.0002293616,0.000573666,0.9193994,0.0064992793,0.04552826,0.024846151,0.00008255341],"about_ca_topic_score_codex":0.006766238,"about_ca_topic_score_gemma":0.007231565,"teacher_disagreement_score":0.006766238,"about_ca_system_score_codex":0.00090837764,"about_ca_system_score_gemma":0.0011000899,"threshold_uncertainty_score":0.013703704},"labels":[],"label_agreement":null},{"id":"W2250867360","doi":"10.3115/v1/p15-4008","title":"A Web-based Collaborative Evaluation Tool for Automatically Learned Relation Extraction Patterns","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Computer science; Preprocessor; Relationship extraction; Dependency (UML); Parsing; Artificial intelligence; Annotation; Dependency grammar; Categorization; Natural language processing; Relation (database); Quality (philosophy); Information extraction; Machine learning; Data mining","score_opus":0.07510523306801112,"score_gpt":0.35219912374496926,"score_spread":0.2770938906769581,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250867360","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10543779,0.00060407206,0.6187368,0.0005260696,0.00030186799,0.0022857492,0.02169444,0.23860547,0.011807836],"genre_scores_gemma":[0.3119866,0.00021075814,0.60041565,0.0003057087,0.00014291581,0.0041473443,0.06283367,0.01091571,0.00904163],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98521113,0.005803699,0.0016390863,0.0027948923,0.0041339668,0.0004171381],"domain_scores_gemma":[0.931988,0.04134568,0.0025364321,0.0110352095,0.011410485,0.0016841311],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014406486,0.0030388823,0.001995702,0.00940634,0.0014929639,0.0027810112,0.0033217317,0.00238656,0.01333603],"category_scores_gemma":[0.056926716,0.00091027765,0.001135258,0.004936875,0.00059610006,0.005618853,0.0032707867,0.0018744023,0.0067031905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001920988,0.0028963734,0.016528705,0.0018064593,0.0007502085,0.0010315472,0.0019148533,0.014790674,0.033768024,0.004232906,0.15253463,0.7678247],"study_design_scores_gemma":[0.0009816389,0.0015448208,0.028086975,0.0002227188,0.00032355613,0.0013087277,0.0010669653,0.8039562,0.06538662,0.012940567,0.08379441,0.00038678868],"about_ca_topic_score_codex":0.004592575,"about_ca_topic_score_gemma":0.008001552,"teacher_disagreement_score":0.014406486,"about_ca_system_score_codex":0.001153259,"about_ca_system_score_gemma":0.0017923428,"threshold_uncertainty_score":0.0761897},"labels":[],"label_agreement":null},{"id":"W2250974997","doi":"10.63317/29jpe4q7e5co","title":"Improving Open Relation Extraction via Sentence Re-Structuring","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Relationship extraction; Structuring; Natural language processing; Information extraction; Artificial intelligence; Predicate (mathematical logic); Sentence; Natural language; Open domain; Exploit; Natural language understanding; Relation (database); Context (archaeology); Information retrieval; Data mining; Question answering; Programming language","score_opus":0.025038183070540183,"score_gpt":0.26525676000726467,"score_spread":0.2402185769367245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250974997","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04234566,0.001559088,0.9378889,0.00056124624,0.00029327156,0.0005756978,0.0017738,0.0120252855,0.0029771228],"genre_scores_gemma":[0.09648589,0.0006814173,0.8898863,0.00025250725,0.0002980687,0.00025697236,0.007538787,0.0009265348,0.003673535],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99772424,0.0007197376,0.00029231343,0.00058451056,0.00058674236,0.00009244491],"domain_scores_gemma":[0.99054986,0.005498899,0.0006902043,0.001296907,0.0018472782,0.00011687121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021354754,0.0017682525,0.0014844007,0.003062342,0.0008772062,0.0016018396,0.0014037323,0.0009646891,0.004146322],"category_scores_gemma":[0.009723631,0.00053325883,0.0014014662,0.0025364792,0.00051148044,0.0040665288,0.0020453911,0.0017631869,0.0044568195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031324915,0.0002948177,0.0019595276,0.0008738959,0.00014608682,0.0006034907,0.0015151952,0.0051504206,0.11374316,0.005776119,0.019348882,0.85027516],"study_design_scores_gemma":[0.00018953753,0.0007503393,0.0089329425,0.00028081078,0.00094573695,0.0023939507,0.00236798,0.4650305,0.31686702,0.041722793,0.16022764,0.00029071968],"about_ca_topic_score_codex":0.0012275664,"about_ca_topic_score_gemma":0.0020773911,"teacher_disagreement_score":0.004146322,"about_ca_system_score_codex":0.00036619222,"about_ca_system_score_gemma":0.0012773199,"threshold_uncertainty_score":0.013870776},"labels":[],"label_agreement":null},{"id":"W2251141757","doi":"10.3115/v1/s14-1007","title":"Identifying semantic relations in a specialized corpus through distributional analysis of a cooccurrence tensor","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Natural language processing; Tensor (intrinsic definition); Encoding (memory); Semantic space; Artificial intelligence; Space (punctuation); Word (group theory); Style (visual arts); Linguistics; Mathematics","score_opus":0.04415604225102656,"score_gpt":0.3008423448311019,"score_spread":0.25668630258007535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251141757","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038891032,0.00017423394,0.95684177,0.00019019922,0.000034576664,0.0000989312,0.00075978285,0.000929569,0.002079817],"genre_scores_gemma":[0.47246304,0.0002915545,0.52117765,0.000086821885,0.000070084425,0.0003506321,0.0025113334,0.00037268733,0.0026762784],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980667,0.0006268993,0.00016547712,0.0005736683,0.00042894407,0.00013832263],"domain_scores_gemma":[0.9940118,0.002620216,0.00066799956,0.0013286031,0.0010988744,0.0002724453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002471446,0.0007853878,0.00064105703,0.0048154504,0.0012638578,0.0023603456,0.001018926,0.0006920066,0.003496286],"category_scores_gemma":[0.012642698,0.00040388256,0.001081842,0.004966273,0.0014786568,0.004434342,0.0022670396,0.0017938116,0.0013701649],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047063714,0.00030777513,0.027247211,0.0007466091,0.00027433963,0.0006544908,0.005544242,0.035582863,0.08297155,0.21528198,0.009810703,0.6211076],"study_design_scores_gemma":[0.000019583695,0.00013062826,0.012592697,0.00008470744,0.00010813632,0.0006654371,0.001571879,0.7073005,0.019263523,0.2396266,0.01849958,0.00013672671],"about_ca_topic_score_codex":0.0071227467,"about_ca_topic_score_gemma":0.012463355,"teacher_disagreement_score":0.0071227467,"about_ca_system_score_codex":0.0012302862,"about_ca_system_score_gemma":0.0018377142,"threshold_uncertainty_score":0.0141626},"labels":[],"label_agreement":null},{"id":"W2251363724","doi":"10.3115/v1/p15-1075","title":"Efficient Methods for Inferring Large Sparse Topic Hierarchies","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Science North","funders":"Defense Advanced Research Projects Agency; Allen Institute for Artificial Intelligence; National Science Foundation","keywords":"Computer science; Computational linguistics; Joint (building); Volume (thermodynamics); Natural language processing; Artificial intelligence; Data science; Engineering; Physics","score_opus":0.11025631062851733,"score_gpt":0.382636016183524,"score_spread":0.2723797055550067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251363724","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008407648,0.0020192147,0.9841905,0.0004613729,0.00009180095,0.000120368364,0.0010828464,0.0025222944,0.001104033],"genre_scores_gemma":[0.20715858,0.0016791034,0.77494985,0.00032520856,0.0006299023,0.0005109544,0.0100667365,0.000770531,0.0039090863],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99753344,0.001010636,0.00014400497,0.0005959251,0.00050002016,0.00021588575],"domain_scores_gemma":[0.9911608,0.0065780045,0.00032762135,0.000954158,0.00075747684,0.00022194626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031134423,0.0015261424,0.0018275367,0.0046848855,0.0015781305,0.0025660098,0.0032756424,0.0018892455,0.0036623564],"category_scores_gemma":[0.016430382,0.001583983,0.0021012793,0.005387316,0.0007553541,0.0048358333,0.0029060994,0.003340031,0.002589702],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006223901,0.0003682496,0.00877146,0.0007057468,0.00069994014,0.00030941606,0.0014271291,0.102925554,0.006989487,0.055065617,0.07412387,0.74799126],"study_design_scores_gemma":[0.00010434932,0.000030200246,0.0010660569,0.000050108763,0.00011825149,0.00014120535,0.0002055394,0.86613065,0.0013079596,0.122818165,0.00799918,0.000028260765],"about_ca_topic_score_codex":0.013324,"about_ca_topic_score_gemma":0.032709565,"teacher_disagreement_score":0.013324,"about_ca_system_score_codex":0.0010891188,"about_ca_system_score_gemma":0.0016581141,"threshold_uncertainty_score":0.026492894},"labels":[],"label_agreement":null},{"id":"W2251564431","doi":"","title":"On the Effectiveness of using Sentence Compression Models for Query-Focused Multi-Document Summarization","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Redundancy (engineering); Artificial intelligence; Cosine similarity; Natural language processing; Sentence; Relevance (law); Semantic similarity; Pattern recognition (psychology)","score_opus":0.07576314842251103,"score_gpt":0.30272608976452753,"score_spread":0.2269629413420165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251564431","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17691803,0.0072815306,0.80594987,0.0014462718,0.00015607213,0.00034350556,0.00046279523,0.0027359936,0.0047058878],"genre_scores_gemma":[0.7621002,0.0027113382,0.23014644,0.00030790558,0.00029472265,0.0002301861,0.0015403016,0.0002311818,0.0024376586],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99845374,0.00084778754,0.00008882716,0.00024628887,0.0002854996,0.00007779843],"domain_scores_gemma":[0.99157035,0.006780639,0.0003536687,0.0005301693,0.0006676119,0.000097567696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033326042,0.0013532696,0.0013537981,0.0011677403,0.0004132131,0.0013163666,0.0009295762,0.00110268,0.0013577259],"category_scores_gemma":[0.011365995,0.00033130957,0.0007711602,0.0012011942,0.0004735861,0.0031299356,0.0006099426,0.0009751795,0.0005905536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009880741,0.00037760707,0.0029305364,0.0005629887,0.00033857938,0.00013612586,0.0002411091,0.53833777,0.02191386,0.0065954453,0.0037258563,0.42385206],"study_design_scores_gemma":[0.00003284167,0.00031534306,0.00066398777,0.000018999292,0.00008455633,0.00004676583,0.00003729973,0.989225,0.006501201,0.0023163098,0.0007384143,0.000019293206],"about_ca_topic_score_codex":0.003173823,"about_ca_topic_score_gemma":0.0037209147,"teacher_disagreement_score":0.0033326042,"about_ca_system_score_codex":0.0008091469,"about_ca_system_score_gemma":0.0008632783,"threshold_uncertainty_score":0.017624676},"labels":[],"label_agreement":null},{"id":"W2251640092","doi":"10.3115/v1/p15-2081","title":"The Fixed-Size Ordinally-Forgetting Encoding Method for Neural Network Language Models","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Fundamental Research Funds for the Central Universities; Natural Sciences and Engineering Research Council of Canada; Australian Government","keywords":"Forgetting; Encoding (memory); Zhàng; Computer science; Artificial neural network; Computational linguistics; Natural language processing; Artificial intelligence; Natural language; Cognitive science; Linguistics; Philosophy; China; Psychology; Political science; Law","score_opus":0.05572273878499998,"score_gpt":0.31588203900422707,"score_spread":0.26015930021922706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251640092","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009535673,0.0005592664,0.9861064,0.00018951256,0.00015913164,0.000033742712,0.00023682906,0.0024391448,0.00074023096],"genre_scores_gemma":[0.3568902,0.0007865675,0.63437754,0.00022472482,0.00030187907,0.0002933452,0.0019557565,0.0006942629,0.004475636],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990748,0.00033019256,0.00011254235,0.00020605326,0.00018107337,0.00009530837],"domain_scores_gemma":[0.9967803,0.0014831045,0.00014768224,0.000857171,0.0006116935,0.00011998903],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002033906,0.0010532913,0.0011610349,0.0012735788,0.000569796,0.0012242775,0.0031745892,0.0009450966,0.005227093],"category_scores_gemma":[0.00954783,0.0005963933,0.0010761506,0.0012382452,0.0005462978,0.004971318,0.0014002182,0.0030811573,0.001554731],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004849987,0.00020962775,0.001279034,0.00020957178,0.000109486056,0.00013050836,0.00014056584,0.13897434,0.003608085,0.045408312,0.009519851,0.7999256],"study_design_scores_gemma":[0.000029117973,0.000035755653,0.000111181085,0.000015764499,0.00003191832,0.000027868635,0.000011072265,0.9682769,0.0016817865,0.028540464,0.0012242051,0.0000138476325],"about_ca_topic_score_codex":0.005503229,"about_ca_topic_score_gemma":0.007697636,"teacher_disagreement_score":0.005503229,"about_ca_system_score_codex":0.00095107657,"about_ca_system_score_gemma":0.0016043145,"threshold_uncertainty_score":0.017486334},"labels":[],"label_agreement":null},{"id":"W2251693562","doi":"10.3115/v1/p15-2025","title":"Representation Based Translation Evaluation Metrics","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Chen; Computer science; Natural language processing; Representation (politics); Computational linguistics; Volume (thermodynamics); Joint (building); Translation (biology); Artificial intelligence; Association (psychology); Machine translation; Linguistics; Programming language; Engineering; Political science","score_opus":0.3391495938485538,"score_gpt":0.388191459149994,"score_spread":0.04904186530144017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251693562","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30080634,0.042235617,0.573213,0.0040134713,0.0026951095,0.0014598748,0.01415264,0.010545775,0.050878122],"genre_scores_gemma":[0.7477977,0.0032892225,0.2071736,0.00032897445,0.00045369478,0.0008037224,0.03182495,0.0013653751,0.006962813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9732208,0.014273658,0.0022644335,0.0018327199,0.007467509,0.0009408561],"domain_scores_gemma":[0.9796178,0.008426024,0.0010459613,0.0032158843,0.0071216077,0.00057281926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0125015145,0.001722331,0.0018994594,0.010470924,0.0012530253,0.0038882438,0.0016527136,0.0019216512,0.0043411246],"category_scores_gemma":[0.042016532,0.0003496502,0.0014628912,0.010270104,0.0008509256,0.0049991244,0.0027775737,0.0015258729,0.0022736823],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001576044,0.0006343924,0.012191415,0.0016502724,0.0008149924,0.00022271878,0.00043677643,0.044168826,0.012213036,0.028854478,0.06316378,0.83407336],"study_design_scores_gemma":[0.00042117917,0.00298583,0.020736977,0.0007225154,0.0011748811,0.0013529535,0.0008883007,0.8003595,0.041341517,0.076363996,0.053331804,0.0003205473],"about_ca_topic_score_codex":0.0031033827,"about_ca_topic_score_gemma":0.0033368578,"teacher_disagreement_score":0.0125015145,"about_ca_system_score_codex":0.0024822603,"about_ca_system_score_gemma":0.0020918616,"threshold_uncertainty_score":0.06611508},"labels":[],"label_agreement":null},{"id":"W2251942550","doi":"10.3115/v1/p15-1051","title":"Encoding Distributional Semantics into Triple-Based Knowledge Ranking for Document Enrichment","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Semantics (computer science); Computer science; Computational linguistics; Zhàng; Ranking (information retrieval); Natural language processing; Encoding (memory); Volume (thermodynamics); Artificial intelligence; Linguistics; Association (psychology); Information retrieval; Library science; Programming language; History; China; Philosophy; Epistemology; Archaeology","score_opus":0.04827935784633595,"score_gpt":0.3117980470521634,"score_spread":0.26351868920582744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251942550","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07871309,0.0023522929,0.896285,0.00072212005,0.00037228744,0.00036005306,0.007679301,0.008310628,0.005205295],"genre_scores_gemma":[0.47310144,0.0012230619,0.5014382,0.00026055393,0.00019658232,0.00038101006,0.019411772,0.00049533293,0.0034920594],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991227,0.00017623916,0.00012323205,0.00021707102,0.00025464236,0.00010607292],"domain_scores_gemma":[0.9978283,0.00080653257,0.00011588097,0.00038632329,0.0007393311,0.00012367833],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008921722,0.000744341,0.0012898498,0.005840595,0.0007943053,0.0016675433,0.001352007,0.00076708075,0.0030195136],"category_scores_gemma":[0.0044351965,0.00031874835,0.0011844423,0.006583501,0.00035477508,0.0048473347,0.0021649501,0.0009935295,0.0015612689],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005990384,0.0005879766,0.00506845,0.0005118719,0.00022262338,0.00027783425,0.00039229234,0.027095193,0.0145634925,0.027110657,0.034797978,0.8887726],"study_design_scores_gemma":[0.00010998831,0.00027909924,0.0023338427,0.00013016147,0.00026219574,0.00027118617,0.0005496773,0.8213655,0.015862266,0.14223027,0.016512895,0.00009299888],"about_ca_topic_score_codex":0.0054781423,"about_ca_topic_score_gemma":0.01542275,"teacher_disagreement_score":0.005840595,"about_ca_system_score_codex":0.00092616293,"about_ca_system_score_gemma":0.0016670418,"threshold_uncertainty_score":0.01089251},"labels":[],"label_agreement":null},{"id":"W2251998605","doi":"","title":"Hybrid Models for Lexical Acquisition of Correlated Styles","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Lexicon; Computer science; Natural language processing; Artificial intelligence; Variety (cybernetics); Focus (optics); Task (project management); Sentiment analysis; Distributional semantics; Style (visual arts); Machine learning; Semantic similarity","score_opus":0.02671839369732189,"score_gpt":0.2402402779618865,"score_spread":0.2135218842645646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251998605","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029457802,0.00030195704,0.96464545,0.0003920556,0.000049939914,0.00012369304,0.0005679535,0.0013138175,0.0031473155],"genre_scores_gemma":[0.57861453,0.0005998809,0.40430138,0.00028975288,0.00019440272,0.00075846905,0.0026045449,0.0005761924,0.012060917],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988231,0.00050665817,0.000086629574,0.00032245103,0.00017001767,0.00009103046],"domain_scores_gemma":[0.99362564,0.00466821,0.00033838514,0.0006408128,0.0005674494,0.00015939627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003166319,0.0010093137,0.0009971467,0.002680226,0.00069125387,0.003040897,0.0021191905,0.001384724,0.005956205],"category_scores_gemma":[0.011913069,0.0010581334,0.0016334443,0.0022672983,0.0010281898,0.0058224713,0.0019831962,0.0020345943,0.0027979095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005976729,0.00030477493,0.010422596,0.0003169578,0.0004707423,0.00037223406,0.001354663,0.47321326,0.008303134,0.16744615,0.007932459,0.32926545],"study_design_scores_gemma":[0.000016818965,0.00001761003,0.00042386315,0.00001095445,0.000019528994,0.0000414087,0.000032029184,0.95237625,0.00037482084,0.045637712,0.0010293114,0.000019712706],"about_ca_topic_score_codex":0.005198405,"about_ca_topic_score_gemma":0.011273172,"teacher_disagreement_score":0.005956205,"about_ca_system_score_codex":0.0010976311,"about_ca_system_score_gemma":0.00094760547,"threshold_uncertainty_score":0.019925475},"labels":[],"label_agreement":null},{"id":"W2252029281","doi":"","title":"Abstractive Meeting Summarization with Entailment and Fusion","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":76,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of British Columbia","funders":"","keywords":"Computer science; Automatic summarization; Grammaticality; Natural language processing; Artificial intelligence; Sentence; Graph; Exploit; Word (group theory); Textual entailment; Logical consequence; Grammar; Linguistics; Theoretical computer science","score_opus":0.008454136821494991,"score_gpt":0.1986889439120823,"score_spread":0.19023480709058732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252029281","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037205617,0.00020312394,0.9922043,0.00014531573,0.000059890343,0.00013939205,0.0004492483,0.0024001542,0.00067801255],"genre_scores_gemma":[0.07328474,0.0002535797,0.917891,0.00014230421,0.00021438149,0.00034136025,0.004636549,0.000488223,0.002747879],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99654526,0.0010748049,0.00034519698,0.00089057145,0.0009395855,0.00020462583],"domain_scores_gemma":[0.99619883,0.0013336511,0.00041863206,0.00066272565,0.0012416914,0.0001444279],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024955026,0.0027970567,0.0020824703,0.0035381333,0.0013966131,0.0027257109,0.0031485404,0.0020174973,0.0052714283],"category_scores_gemma":[0.008977637,0.00090819766,0.002592774,0.0027225371,0.0007428699,0.00477004,0.0035033587,0.002595038,0.004226264],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006633199,0.00033791925,0.001093211,0.0010756825,0.00042277566,0.0006704582,0.0012391737,0.07137787,0.06109515,0.027200976,0.021642415,0.81318104],"study_design_scores_gemma":[0.000064040665,0.0002745547,0.00082541176,0.00007117567,0.00026816194,0.00025471614,0.00035241287,0.8646032,0.051728424,0.061524957,0.01993251,0.000100433725],"about_ca_topic_score_codex":0.0030353665,"about_ca_topic_score_gemma":0.0045233057,"teacher_disagreement_score":0.0052714283,"about_ca_system_score_codex":0.0009728117,"about_ca_system_score_gemma":0.0015828566,"threshold_uncertainty_score":0.01763469},"labels":[],"label_agreement":null},{"id":"W2252062369","doi":"","title":"CLaC-CORE: Exhaustive Feature Combination for Measuring Textual Similarity","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Feature (linguistics); Computer science; Similarity (geometry); Feature vector; Artificial intelligence; Set (abstract data type); Task (project management); Train; Core (optical fiber); Pattern recognition (psychology); Vector space model; Natural language processing; Data mining; Engineering","score_opus":0.05526690881528108,"score_gpt":0.2600293272245777,"score_spread":0.20476241840929665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252062369","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36062956,0.0016846943,0.57173216,0.00023397896,0.00026733114,0.002117086,0.021788675,0.03299319,0.0085533485],"genre_scores_gemma":[0.6424518,0.00018008916,0.3154118,0.00012754083,0.000092055045,0.0018137011,0.035536025,0.00081282115,0.003574129],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99664503,0.00076664425,0.00034357185,0.00067331456,0.0013010531,0.00027028337],"domain_scores_gemma":[0.9937086,0.0027096928,0.0005306142,0.00091748225,0.0017864872,0.0003471302],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003365605,0.0015615186,0.0014541751,0.0072820866,0.0007850294,0.0012262628,0.001591229,0.0015233315,0.004244071],"category_scores_gemma":[0.012965559,0.0003777796,0.0009432853,0.0038785876,0.00036964443,0.0027805048,0.0027132456,0.00082260853,0.0026291758],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017754402,0.0010344014,0.044314694,0.0011456071,0.00085622,0.00033087193,0.0007562521,0.010598595,0.05203727,0.002025275,0.048978228,0.83614707],"study_design_scores_gemma":[0.00037700418,0.0022826577,0.13480972,0.00014355924,0.0006548967,0.0015932928,0.0008287115,0.7550195,0.067647286,0.011112125,0.025135027,0.00039633331],"about_ca_topic_score_codex":0.0033616838,"about_ca_topic_score_gemma":0.005356275,"teacher_disagreement_score":0.0072820866,"about_ca_system_score_codex":0.00064773124,"about_ca_system_score_gemma":0.0011074535,"threshold_uncertainty_score":0.017799258},"labels":[],"label_agreement":null},{"id":"W2252143850","doi":"","title":"Training recurrent neural networks","year":2013,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":391,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Recurrent neural network; Initialization; Computer science; Sequence (biology); Artificial intelligence; Probabilistic logic; Gradient descent; Artificial neural network; Machine learning","score_opus":0.053348200930180216,"score_gpt":0.2900405536168376,"score_spread":0.2366923526866574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252143850","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031293646,0.0005644902,0.9614362,0.00027528915,0.00010226881,0.00005505772,0.0001360819,0.0014516925,0.004685196],"genre_scores_gemma":[0.60974115,0.0008423302,0.37957695,0.00028679252,0.000109561486,0.00026609126,0.0008437416,0.00042924928,0.00790417],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994215,0.00015876826,0.00003974407,0.00018302874,0.0001368809,0.000060005805],"domain_scores_gemma":[0.9988716,0.00062112877,0.00008946331,0.00017227557,0.00021227919,0.000033243472],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012464515,0.0009404979,0.00080545974,0.0004271472,0.00032130972,0.00070972217,0.0012505173,0.0010373559,0.0025236174],"category_scores_gemma":[0.005072165,0.0007340946,0.0007079023,0.00045034243,0.00049123465,0.0016887569,0.0008513911,0.0014757814,0.0010445387],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004351595,0.00004126527,0.00078373286,0.00012805844,0.00007154793,0.000043966258,0.00006913942,0.854232,0.0052855574,0.014102088,0.0020471702,0.12315192],"study_design_scores_gemma":[0.000003198337,0.0000131923925,0.000060705614,0.000007467064,0.0000055264945,0.0000065348077,0.0000037926752,0.9946853,0.00121044,0.0033988818,0.0006016887,0.0000031875213],"about_ca_topic_score_codex":0.0037577318,"about_ca_topic_score_gemma":0.005838322,"teacher_disagreement_score":0.0037577318,"about_ca_system_score_codex":0.00072099327,"about_ca_system_score_gemma":0.0007994781,"threshold_uncertainty_score":0.008442402},"labels":[],"label_agreement":null},{"id":"W2252230762","doi":"","title":"Using Syntactic and Shallow Semantic Kernels to Improve Multi-Modality Manifold-Ranking for Topic-Focused Multi-Document Summarization","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Ranking (information retrieval); Natural language processing; Cosine similarity; Relevance (law); Artificial intelligence; Information retrieval; Benchmark (surveying); Similarity (geometry); Semantic similarity; Modality (human–computer interaction); Pattern recognition (psychology); Image (mathematics)","score_opus":0.11503794682435643,"score_gpt":0.30975177393011927,"score_spread":0.19471382710576285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252230762","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0355908,0.00047502172,0.961202,0.00012756196,0.000042218187,0.000060984417,0.00010852785,0.001830881,0.00056199677],"genre_scores_gemma":[0.53271943,0.0003520269,0.46223715,0.00012516213,0.00016667519,0.00015082554,0.0014588939,0.00028554373,0.0025042098],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989492,0.00043801568,0.00008371883,0.00017482114,0.00025728848,0.000096962],"domain_scores_gemma":[0.99782926,0.000685504,0.00024648858,0.00045033643,0.0006912568,0.00009722189],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016249223,0.0010358716,0.0014766109,0.0022332678,0.0006500522,0.0011327632,0.0010374784,0.0011328949,0.0013511248],"category_scores_gemma":[0.0046751923,0.00029318634,0.00096346706,0.0017902564,0.00050467136,0.0030564577,0.000989732,0.0011018032,0.0011630318],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040120934,0.00048192826,0.002508361,0.0002837358,0.0002795749,0.00017217474,0.0003555657,0.09844514,0.072988,0.0131613165,0.007986335,0.8029367],"study_design_scores_gemma":[0.000022334618,0.00017726862,0.001030102,0.00000839625,0.000052116186,0.000077190234,0.00006667637,0.9757902,0.011352612,0.0099883685,0.0013900618,0.000044691285],"about_ca_topic_score_codex":0.0019842205,"about_ca_topic_score_gemma":0.0040937075,"teacher_disagreement_score":0.0022332678,"about_ca_system_score_codex":0.00058522745,"about_ca_system_score_gemma":0.0007493776,"threshold_uncertainty_score":0.008593559},"labels":[],"label_agreement":null},{"id":"W2252242089","doi":"","title":"On the Effectiveness of Using Syntactic and Shallow Semantic Tree Kernels for Automatic Assessment of Essays","year":2013,"lang":"en","type":"article","venue":"International Joint Conference on Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Natural language processing; Grading (engineering); Artificial intelligence; Latent semantic analysis; Task (project management); Tree (set theory); Tree structure; Data structure; Programming language","score_opus":0.031931837597632816,"score_gpt":0.3179565227045993,"score_spread":0.2860246851069665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252242089","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60745895,0.0022237888,0.38109407,0.0012119131,0.00015506562,0.00011926473,0.00021292365,0.002417459,0.00510664],"genre_scores_gemma":[0.9511302,0.00031211376,0.046720125,0.000079218364,0.00004764641,0.000023620807,0.0002821313,0.00007582788,0.0013290939],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960114,0.0022657567,0.00026205613,0.00057225244,0.0006394,0.00024923027],"domain_scores_gemma":[0.961245,0.03061359,0.0011553549,0.0023454593,0.0039917417,0.0006488019],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009898956,0.0012011811,0.00095631904,0.0018476186,0.0008028311,0.0023761482,0.0010719901,0.0021743928,0.001296794],"category_scores_gemma":[0.03698063,0.00047647973,0.00065865036,0.0009928595,0.00087834033,0.006000955,0.002018352,0.0020258296,0.00093815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027777902,0.0010923271,0.021694468,0.00023536812,0.00030805334,0.00011755108,0.00049930485,0.20308927,0.02212392,0.009571344,0.003492478,0.7349982],"study_design_scores_gemma":[0.000022141185,0.00011178989,0.0022164194,0.000012260452,0.00003473412,0.000023185115,0.000053732034,0.99102634,0.0032434834,0.0030270312,0.00021059767,0.000018194545],"about_ca_topic_score_codex":0.008450241,"about_ca_topic_score_gemma":0.0069870716,"teacher_disagreement_score":0.009898956,"about_ca_system_score_codex":0.0009466435,"about_ca_system_score_gemma":0.0011854555,"threshold_uncertainty_score":0.052351356},"labels":[],"label_agreement":null},{"id":"W2275754865","doi":"","title":"Commentary on van Eemeren & Houtlosser","year":2001,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Epistemology; Philosophy","score_opus":0.027777221172339137,"score_gpt":0.2614952193566611,"score_spread":0.23371799818432196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2275754865","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00003793137,0.00246192,0.00007180434,0.96116096,0.035455238,0.000004868021,0.00007081412,0.000009429321,0.0007271476],"genre_scores_gemma":[0.0011308993,0.0018955887,0.00016776256,0.94759387,0.046171717,0.00004225133,0.000044812623,0.000041564002,0.0029115365],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9832598,0.0061605857,0.0017364405,0.0024193558,0.005256649,0.0011671592],"domain_scores_gemma":[0.8670586,0.103852846,0.0036341073,0.0020304557,0.019230435,0.0041936524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025875587,0.0018018795,0.0031801427,0.003725342,0.008277409,0.009812679,0.007311162,0.067741744,0.011448499],"category_scores_gemma":[0.15626653,0.0013099831,0.0029329807,0.0047444995,0.007065852,0.009036781,0.00445385,0.06901706,0.00857021],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00001655133,0.0000030261683,0.00002337181,0.000064025175,0.000011489606,0.000032269712,0.00015535708,0.000016515552,0.0000125193965,0.0017154888,0.9962226,0.0017266539],"study_design_scores_gemma":[0.00006509594,0.000012985905,0.00041777437,0.0012324434,0.000092615854,0.000121256555,0.0007162903,0.00013198602,0.000094654795,0.008993287,0.98807085,0.00005070605],"about_ca_topic_score_codex":0.043774355,"about_ca_topic_score_gemma":0.059761144,"teacher_disagreement_score":0.067741744,"about_ca_system_score_codex":0.010864471,"about_ca_system_score_gemma":0.011884483,"threshold_uncertainty_score":0.13684481},"labels":[],"label_agreement":null},{"id":"W2279735888","doi":"10.1109/icdmw.2015.108","title":"Comparing SVD and SDAE for Analysis of Islamist Forum Postings","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Singular value decomposition; Data science; Artificial intelligence","score_opus":0.08284314093834615,"score_gpt":0.29039010855105407,"score_spread":0.2075469676127079,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2279735888","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.66743404,0.0014015334,0.32098076,0.0006417658,0.0002668048,0.00011582417,0.0017880324,0.0023918455,0.004979401],"genre_scores_gemma":[0.84466285,0.0005570546,0.14890622,0.00007390001,0.00014527474,0.00007150533,0.0037435975,0.00013298068,0.0017064917],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99853086,0.0005800647,0.00013000809,0.000287498,0.0003248794,0.0001466781],"domain_scores_gemma":[0.98872954,0.008412385,0.00041912275,0.00065494736,0.001433049,0.00035100733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002891134,0.0007018449,0.0006910389,0.0045279157,0.00029855248,0.0014118474,0.00035497994,0.00082450337,0.0014952637],"category_scores_gemma":[0.012203682,0.00021817007,0.00092794397,0.002350561,0.0003535525,0.0020650309,0.0007856285,0.0010549597,0.0010913647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017858709,0.0006386718,0.06467225,0.00055326754,0.00069549173,0.00018359766,0.0013762972,0.07837666,0.03438043,0.0060848435,0.0040655206,0.8071872],"study_design_scores_gemma":[0.00003230548,0.0002679717,0.036051683,0.00004501962,0.00007384298,0.00014878204,0.0009915128,0.943629,0.011044324,0.005657728,0.0019903998,0.000067324625],"about_ca_topic_score_codex":0.0037705167,"about_ca_topic_score_gemma":0.005581622,"teacher_disagreement_score":0.0045279157,"about_ca_system_score_codex":0.00031919172,"about_ca_system_score_gemma":0.00043706788,"threshold_uncertainty_score":0.015289962},"labels":[],"label_agreement":null},{"id":"W2283763484","doi":"10.1145/2651444","title":"Recognition of Patient-Related Named Entities in Noisy Tele-Health Texts","year":2015,"lang":"en","type":"article","venue":"ACM Transactions on Intelligent Systems and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Relationship extraction; Noise (video); Information extraction; Information retrieval; Sentence; Named-entity recognition; Phone; Filter (signal processing); Machine learning; Task (project management); Speech recognition","score_opus":0.03429504103428756,"score_gpt":0.2580997059002944,"score_spread":0.22380466486600686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2283763484","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32677513,0.0016568666,0.652484,0.0025331548,0.00024019377,0.0004659769,0.009139421,0.0030223855,0.003682897],"genre_scores_gemma":[0.52741754,0.0008241617,0.45851427,0.0003631127,0.0002618702,0.00033114012,0.0100140525,0.00027840445,0.001995434],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9967347,0.0011910094,0.00047014945,0.00081382017,0.0006693202,0.000120956654],"domain_scores_gemma":[0.9774418,0.015782045,0.0037429077,0.0013073559,0.0015276582,0.00019822727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026505992,0.00076540135,0.0005560298,0.0032365662,0.00053045165,0.0016177779,0.0010209448,0.0015253404,0.0013035845],"category_scores_gemma":[0.016065415,0.0004215353,0.0004612916,0.0027227083,0.00087069423,0.0029928132,0.0010937565,0.00090800115,0.0010103538],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015036484,0.0005416124,0.070052415,0.0038945808,0.00026445568,0.009370404,0.018980779,0.049423583,0.17723043,0.026497673,0.015931627,0.6263088],"study_design_scores_gemma":[0.00012092923,0.0003979985,0.075536065,0.00057599065,0.00029000922,0.0050613876,0.0068992306,0.62402993,0.13774385,0.058253575,0.09082671,0.0002643023],"about_ca_topic_score_codex":0.0014354317,"about_ca_topic_score_gemma":0.0022331232,"teacher_disagreement_score":0.0032365662,"about_ca_system_score_codex":0.0006943485,"about_ca_system_score_gemma":0.0007549339,"threshold_uncertainty_score":0.01401788},"labels":[],"label_agreement":null},{"id":"W2289035630","doi":"10.1109/asru.2015.7404821","title":"Recent improvements to NeuroCRFs for named entity recognition","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Conditional random field; Artificial intelligence; Margin (machine learning); Computer science; Pattern recognition (psychology); Sequence labeling; Feature (linguistics); Artificial neural network; Sequence (biology); Task (project management); Component (thermodynamics); Machine learning; Natural language processing","score_opus":0.11358371001075171,"score_gpt":0.29803458051047177,"score_spread":0.18445087049972006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2289035630","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017530043,0.0096185915,0.87714696,0.0021657282,0.0012038758,0.0003900137,0.007761047,0.06599964,0.018184071],"genre_scores_gemma":[0.13209142,0.0045468197,0.7944859,0.0016583776,0.00094662147,0.0006449506,0.03054209,0.002749107,0.032334697],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99467754,0.0010195806,0.00031934996,0.0016202177,0.002040781,0.00032258293],"domain_scores_gemma":[0.9908138,0.002802156,0.0002673998,0.0032070733,0.0026752327,0.0002343077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061046807,0.0028980523,0.0016459659,0.004462883,0.0014525807,0.0023698977,0.0052962047,0.0022318996,0.012621109],"category_scores_gemma":[0.01399259,0.0010378244,0.002016004,0.004817107,0.00083062745,0.00686658,0.0024243987,0.003114887,0.015326548],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034148875,0.00026959178,0.0019640715,0.00032757933,0.00022205345,0.00009870133,0.00009109389,0.043718234,0.0038604795,0.007949405,0.07463138,0.86652595],"study_design_scores_gemma":[0.0000637856,0.00014039534,0.0023205641,0.00013211052,0.00016179404,0.00044429084,0.0001007366,0.7802622,0.019730395,0.024933042,0.17158487,0.00012591919],"about_ca_topic_score_codex":0.05639606,"about_ca_topic_score_gemma":0.06903235,"teacher_disagreement_score":0.05639606,"about_ca_system_score_codex":0.0031757876,"about_ca_system_score_gemma":0.0037797776,"threshold_uncertainty_score":0.11213559},"labels":[],"label_agreement":null},{"id":"W2289084663","doi":"10.13140/2.1.4822.8166","title":"Towards News Verification: Deception Detection Methods for News Discourse","year":2015,"lang":"en","type":"article","venue":"Scholarship@Western (Western University)","topic":"Topic Modeling","field":"Computer Science","cited_by":133,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Deception; Rhetorical question; Fake news; Computer science; Coherence (philosophical gambling strategy); Similarity (geometry); Feature (linguistics); Natural language processing; Sample (material); Artificial intelligence; Linguistics; Psychology; Social psychology; Mathematics; Internet privacy","score_opus":0.21550922589119192,"score_gpt":0.40500182832543585,"score_spread":0.18949260243424393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2289084663","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029264404,0.00040052063,0.9656412,0.00066210533,0.00005390643,0.00019466678,0.00027846431,0.002317681,0.0011870346],"genre_scores_gemma":[0.35948917,0.0003213288,0.63547224,0.00015318493,0.00019601548,0.00030274742,0.0010505874,0.00038097033,0.0026338368],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99306846,0.004482752,0.00035851894,0.0010095288,0.0008182647,0.00026259228],"domain_scores_gemma":[0.97601014,0.015912494,0.0018261947,0.0026869036,0.0031827716,0.0003814745],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01185116,0.0012512713,0.0010084714,0.0052888594,0.0013708953,0.0044490704,0.0025518734,0.0019513597,0.0030924222],"category_scores_gemma":[0.033834495,0.0007332357,0.0017490134,0.0020735827,0.0013158664,0.0056454116,0.0020414519,0.0032798795,0.002909694],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008143459,0.0005651584,0.023357203,0.0007058119,0.0002935542,0.00028550508,0.0047341525,0.043229338,0.02496682,0.04477807,0.010328917,0.8459412],"study_design_scores_gemma":[0.000029754901,0.00007160547,0.0030032664,0.00006798116,0.000039914165,0.0001320806,0.00071755354,0.95738715,0.009152249,0.025363406,0.003996622,0.000038526716],"about_ca_topic_score_codex":0.0038715987,"about_ca_topic_score_gemma":0.0033669071,"teacher_disagreement_score":0.01185116,"about_ca_system_score_codex":0.001451451,"about_ca_system_score_gemma":0.0015719346,"threshold_uncertainty_score":0.062675655},"labels":[],"label_agreement":null},{"id":"W2290030465","doi":"","title":"Multi-Document Summarization from First Principles","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Automatic summarization; Computer science; Multi-document summarization; Redundancy (engineering); Baseline (sea); Information retrieval; Artificial intelligence","score_opus":0.028421352550889126,"score_gpt":0.24622392614444824,"score_spread":0.21780257359355912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2290030465","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015203954,0.0020756524,0.9434308,0.00067090895,0.00022836389,0.00039841348,0.0023903647,0.030801116,0.0048004594],"genre_scores_gemma":[0.115876146,0.0010316956,0.86329216,0.00026446528,0.0003966461,0.0002660461,0.008077571,0.0010652331,0.009730002],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99860567,0.0003300136,0.00015600801,0.00034845166,0.00049639354,0.000063399115],"domain_scores_gemma":[0.9973271,0.0008453414,0.0003049106,0.000492047,0.0009351028,0.000095469026],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017776061,0.0012224424,0.0010358538,0.0034789843,0.0007808886,0.002232372,0.0013622494,0.0006155992,0.006376044],"category_scores_gemma":[0.0036255645,0.00046879504,0.0011005575,0.0027021372,0.0003181251,0.0023408395,0.0011387884,0.0013472275,0.0045789583],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027583484,0.00014786048,0.00094616896,0.00091677735,0.00020598523,0.00016826781,0.00039905048,0.009124775,0.057025556,0.0056693526,0.038089763,0.88703066],"study_design_scores_gemma":[0.00022924977,0.0009951777,0.00477727,0.00026001965,0.00053823553,0.0010650912,0.00053341326,0.50254244,0.24230617,0.04073424,0.20582937,0.0001894301],"about_ca_topic_score_codex":0.0014904112,"about_ca_topic_score_gemma":0.002808394,"teacher_disagreement_score":0.006376044,"about_ca_system_score_codex":0.0006275695,"about_ca_system_score_gemma":0.0008813668,"threshold_uncertainty_score":0.021330059},"labels":[],"label_agreement":null},{"id":"W2293026843","doi":"10.3115/v1/n15-1115","title":"Encoding World Knowledge in the Evaluation of Local Coherence","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Zhàng; Coherence (philosophical gambling strategy); Encoding (memory); Computer science; Computational linguistics; Linguistics; Artificial intelligence; Natural language processing; Cognitive science; China; History; Philosophy; Psychology; Physics; Quantum mechanics","score_opus":0.22912596399578897,"score_gpt":0.36988553323863904,"score_spread":0.14075956924285007,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2293026843","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5705694,0.0035453327,0.4038725,0.0018993266,0.00022221911,0.00021227806,0.001200215,0.0022660254,0.01621268],"genre_scores_gemma":[0.976223,0.00019223068,0.021967424,0.00006107852,0.000044820143,0.00004923694,0.0006863825,0.0001526099,0.0006232754],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971655,0.0014162868,0.00019958803,0.0005801232,0.00038563323,0.00025286872],"domain_scores_gemma":[0.9769701,0.016699597,0.0011924062,0.0020773343,0.0024605575,0.00060008746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042691687,0.0006912562,0.0009207061,0.002261636,0.000797811,0.0026703747,0.0009075657,0.0010215192,0.0034505432],"category_scores_gemma":[0.03640959,0.00033841614,0.00033694028,0.002028271,0.0011982534,0.008650882,0.0030703843,0.0011574433,0.0004990912],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0045739254,0.00052342255,0.03713388,0.00077333796,0.0004346441,0.00040730074,0.0027547798,0.0958506,0.018567063,0.08297243,0.014559463,0.7414493],"study_design_scores_gemma":[0.00019103228,0.0005130428,0.011480907,0.0001289324,0.00027424694,0.000107663574,0.0011994629,0.847009,0.014168734,0.121708885,0.0031340856,0.00008398885],"about_ca_topic_score_codex":0.0040496504,"about_ca_topic_score_gemma":0.006546003,"teacher_disagreement_score":0.0042691687,"about_ca_system_score_codex":0.0010807032,"about_ca_system_score_gemma":0.0010994509,"threshold_uncertainty_score":0.022577763},"labels":[],"label_agreement":null},{"id":"W2293834552","doi":"10.3115/v1/n15-1099","title":"A Word Embedding Approach to Predicting the Compositionality of Multiword Expressions","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":99,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"Australian Research Council; Australian Government","keywords":"Principle of compositionality; Computer science; Natural language processing; Word embedding; Artificial intelligence; Embedding; Word (group theory); Linguistics","score_opus":0.08948264642058708,"score_gpt":0.31564365763351504,"score_spread":0.22616101121292798,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2293834552","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1591994,0.0026714539,0.8284865,0.00041994517,0.00036841654,0.00026790454,0.002269174,0.0027057463,0.0036113865],"genre_scores_gemma":[0.5917516,0.0013858678,0.39527193,0.00015982807,0.00027089164,0.0003538343,0.00632068,0.0003372626,0.0041481024],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909663,0.00029706548,0.000101739985,0.0002801872,0.00016544168,0.000058951526],"domain_scores_gemma":[0.9975032,0.0011237422,0.00030794222,0.00037315773,0.0005868614,0.00010505211],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011081436,0.0014263762,0.00062645203,0.0033027614,0.00048250557,0.0010749615,0.00070289517,0.00092309667,0.0021584518],"category_scores_gemma":[0.0057529705,0.00030702862,0.0009777187,0.0028702058,0.0004405149,0.004481031,0.0013026177,0.0014965507,0.0016070909],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057234213,0.00044939268,0.016232012,0.0008002159,0.00034708894,0.00027940256,0.0008293414,0.017985977,0.03948967,0.012523906,0.007636482,0.9028542],"study_design_scores_gemma":[0.0000629796,0.00061480945,0.011024315,0.00013373877,0.00030895817,0.0009818277,0.0008533575,0.9027481,0.020371843,0.04547929,0.017282322,0.00013837918],"about_ca_topic_score_codex":0.0013272922,"about_ca_topic_score_gemma":0.0021544746,"teacher_disagreement_score":0.0033027614,"about_ca_system_score_codex":0.00034867457,"about_ca_system_score_gemma":0.0005681956,"threshold_uncertainty_score":0.0072208047},"labels":[],"label_agreement":null},{"id":"W2294102067","doi":"","title":"A Description of ZZ_INFO_TECH System at KBP 2013","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Ranking (information retrieval); Heuristic; Task (project management); Context (archaeology); Similarity (geometry); Filter (signal processing); Support vector machine; Data mining; Information retrieval; Population; Train; Knowledge base; Simple (philosophy); Artificial intelligence; Machine learning; Geography; Engineering","score_opus":0.012717684007598021,"score_gpt":0.21664126023698088,"score_spread":0.20392357622938287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2294102067","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011352474,0.0013009334,0.21229057,0.0010528893,0.00031680497,0.0019881534,0.18389553,0.52961665,0.058185898],"genre_scores_gemma":[0.08478981,0.0010051836,0.19969107,0.0015285487,0.00022363864,0.0031863567,0.63004047,0.02804652,0.051488433],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982571,0.00021905123,0.00022561273,0.0004483703,0.00069097895,0.00015897912],"domain_scores_gemma":[0.99884063,0.00020571989,0.000058583104,0.00033276746,0.00045565105,0.00010669086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013991485,0.0015298299,0.0013080051,0.0033922263,0.0012767268,0.0030608082,0.0032902027,0.0012264906,0.0948707],"category_scores_gemma":[0.0033955197,0.0012069041,0.00075266534,0.0036152778,0.00036005778,0.003566245,0.0020824703,0.0012475947,0.08916218],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009923334,0.00029160862,0.003156727,0.0014180792,0.00016088416,0.0007281648,0.00041015464,0.0035671166,0.026506228,0.004974258,0.77305794,0.1847365],"study_design_scores_gemma":[0.00043759428,0.00026708192,0.009659549,0.00016973216,0.00019015736,0.001551348,0.0003155066,0.047570407,0.04174051,0.0072125574,0.8905718,0.00031371447],"about_ca_topic_score_codex":0.012598324,"about_ca_topic_score_gemma":0.0071626636,"teacher_disagreement_score":0.0948707,"about_ca_system_score_codex":0.0013636189,"about_ca_system_score_gemma":0.0015687317,"threshold_uncertainty_score":0.3173741},"labels":[],"label_agreement":null},{"id":"W2294150579","doi":"","title":"University of Amsterdam at TAC 2011: English slot filling task (Draft)","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Relation (database); Task (project management); Matching (statistics); Tuple; Population; Track (disk drive); Information retrieval; World Wide Web; Artificial intelligence; Database; Mathematics; Statistics; Engineering","score_opus":0.013156579505681413,"score_gpt":0.197969476123265,"score_spread":0.18481289661758357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2294150579","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064363584,0.002247909,0.061125156,0.009203936,0.0052987,0.002642518,0.6063235,0.06630358,0.18249114],"genre_scores_gemma":[0.082980335,0.00042701734,0.06320888,0.0009582616,0.00046312105,0.0023659302,0.75907797,0.008731717,0.08178684],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.996135,0.0011324886,0.00038717425,0.00092119724,0.0008958317,0.0005283698],"domain_scores_gemma":[0.9897259,0.0019441091,0.00023028407,0.0020994027,0.0045641847,0.0014359988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005162245,0.0014445886,0.0017275417,0.001775935,0.0025836073,0.0046158973,0.0020288292,0.0024491851,0.10198764],"category_scores_gemma":[0.01685652,0.0013499253,0.0008319467,0.003144342,0.00044583273,0.003976077,0.0022705325,0.0027558135,0.118736535],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004004086,0.00023603055,0.0008178788,0.0003126279,0.000031253385,0.00019733375,0.00054920866,0.00082647195,0.0025245484,0.00093818514,0.9543703,0.038795713],"study_design_scores_gemma":[0.00040347927,0.00024225329,0.009825837,0.00023734257,0.00006336512,0.00042756856,0.0015113981,0.011669994,0.010538819,0.004381437,0.96050864,0.00018983657],"about_ca_topic_score_codex":0.030608645,"about_ca_topic_score_gemma":0.03906704,"teacher_disagreement_score":0.10198764,"about_ca_system_score_codex":0.0013937391,"about_ca_system_score_gemma":0.003967605,"threshold_uncertainty_score":0.3411826},"labels":[],"label_agreement":null},{"id":"W2294485447","doi":"","title":"Entity Recognition and Linking on Tweets with Random Walks","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Entity linking; Exploit; Named-entity recognition; Natural language processing; Robustness (evolution); Knowledge base; Artificial intelligence; Task (project management); Component (thermodynamics); Random walk; Graph; Information retrieval; Theoretical computer science; Mathematics","score_opus":0.0735964523235494,"score_gpt":0.2417138755869833,"score_spread":0.16811742326343387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2294485447","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07624199,0.0011025706,0.8215245,0.00070324476,0.00041992287,0.00056614843,0.0152198225,0.075710036,0.008511835],"genre_scores_gemma":[0.312238,0.0006248087,0.61553925,0.00042118013,0.00036066526,0.0004812196,0.05021158,0.0022818986,0.017841289],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982559,0.00038735193,0.0001539032,0.00077218754,0.00029709502,0.00013347581],"domain_scores_gemma":[0.9975834,0.0011125148,0.0002004916,0.00074985396,0.00028723147,0.00006656563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012486968,0.0020002895,0.0012188752,0.005433085,0.0012724889,0.0014710068,0.0021049508,0.0019242172,0.0053461255],"category_scores_gemma":[0.0057622977,0.000724994,0.0018138507,0.0044950456,0.00043523815,0.00458653,0.0021360097,0.0013878623,0.008667488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008518015,0.0005194073,0.0114649255,0.0005840259,0.00043909161,0.0011168681,0.00046979255,0.061369523,0.029474117,0.00884836,0.07396685,0.8108952],"study_design_scores_gemma":[0.000039695413,0.00010050637,0.0030617334,0.000048562277,0.000098335484,0.0005598241,0.00015740197,0.9184339,0.032151747,0.019739978,0.025545416,0.00006293542],"about_ca_topic_score_codex":0.0040554493,"about_ca_topic_score_gemma":0.0091685485,"teacher_disagreement_score":0.005433085,"about_ca_system_score_codex":0.00053171563,"about_ca_system_score_gemma":0.0005192242,"threshold_uncertainty_score":0.017884552},"labels":[],"label_agreement":null},{"id":"W2295082727","doi":"","title":"Stanford's 2013 KBP System","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Consistency (knowledge bases); Inference; Computer science; Component (thermodynamics); Data mining; Information retrieval; Artificial intelligence","score_opus":0.007980749315077551,"score_gpt":0.21882852409483045,"score_spread":0.2108477747797529,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2295082727","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01684314,0.0017686846,0.122125655,0.0029518122,0.0015947358,0.0015311149,0.3629063,0.41744745,0.07283112],"genre_scores_gemma":[0.05897292,0.0007840482,0.20109007,0.0011550703,0.0003365276,0.001878688,0.6920559,0.015733987,0.027992858],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966666,0.00055430084,0.0005300097,0.0007896568,0.0012472044,0.00021225364],"domain_scores_gemma":[0.9929004,0.0016912761,0.00023581614,0.0017914054,0.002916628,0.0004646492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037120976,0.0019870517,0.0019389957,0.00537463,0.0025380081,0.0033674096,0.004192054,0.0019265484,0.08539944],"category_scores_gemma":[0.014613542,0.0015248234,0.0010459054,0.0048014517,0.0006247048,0.0076830364,0.005155904,0.0034007633,0.088029206],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000305337,0.00017862338,0.00074269465,0.0007420351,0.000048400107,0.00023643993,0.0003545924,0.0009315504,0.0033484444,0.0036862353,0.8853605,0.104065135],"study_design_scores_gemma":[0.0003851862,0.00015046021,0.0046303766,0.00030040645,0.00012975888,0.0006966797,0.0006101882,0.026288455,0.0116452295,0.010557173,0.94433933,0.0002668514],"about_ca_topic_score_codex":0.036406305,"about_ca_topic_score_gemma":0.03103588,"teacher_disagreement_score":0.08539944,"about_ca_system_score_codex":0.0018894621,"about_ca_system_score_gemma":0.004181813,"threshold_uncertainty_score":0.2856896},"labels":[],"label_agreement":null},{"id":"W2295840922","doi":"","title":"IIIT Hyderabad in Summarization and Knowledge Base Population at TAC 2011","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Task (project management); Natural language processing; Population; Knowledge base; Word (group theory); Artificial intelligence; Information retrieval; Linguistics","score_opus":0.019018220729069668,"score_gpt":0.23861605475814146,"score_spread":0.2195978340290718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2295840922","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58728486,0.0029819377,0.20392457,0.007725568,0.001517642,0.0021175444,0.038154077,0.047648616,0.108645245],"genre_scores_gemma":[0.6160019,0.00079185143,0.19758931,0.0005618895,0.00037476415,0.000598605,0.06976639,0.002227747,0.11208757],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979685,0.0006258843,0.000116182426,0.00045243578,0.00058797194,0.0002490493],"domain_scores_gemma":[0.99528,0.0010944387,0.00023519021,0.0008339203,0.0017804719,0.0007759646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035661196,0.0004791595,0.0005648478,0.0018922457,0.0013263684,0.002082173,0.0010038471,0.0007195577,0.01004039],"category_scores_gemma":[0.0061070654,0.00030098832,0.0003160266,0.0022953933,0.00025731904,0.0019259974,0.0011043698,0.00094159046,0.006855689],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013858379,0.0008793477,0.013087464,0.000455615,0.00011073795,0.0011127244,0.004126031,0.013774971,0.03876005,0.003701479,0.17078997,0.75181586],"study_design_scores_gemma":[0.0003176199,0.0019648375,0.051048063,0.00011207897,0.00020813878,0.00082301133,0.0038578326,0.10079278,0.13793871,0.0048530176,0.6977644,0.00031954522],"about_ca_topic_score_codex":0.013422533,"about_ca_topic_score_gemma":0.015305902,"teacher_disagreement_score":0.013422533,"about_ca_system_score_codex":0.0014205303,"about_ca_system_score_gemma":0.001898274,"threshold_uncertainty_score":0.03358847},"labels":[],"label_agreement":null},{"id":"W2296086669","doi":"","title":"JU_CSE_TAC: Textual Entailment Recognition System at TAC RTE-6","year":2010,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Textual entailment; Computer science; Natural language processing; Task (project management); Sentence; Novelty; Artificial intelligence; Logical consequence; Set (abstract data type); Similarity (geometry); Programming language","score_opus":0.010880359769517584,"score_gpt":0.23220264957364453,"score_spread":0.22132228980412694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296086669","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19476181,0.0021730333,0.29612792,0.0013063024,0.00089399103,0.0021167658,0.03615301,0.4161275,0.05033971],"genre_scores_gemma":[0.44098222,0.00041243312,0.3689108,0.00081067963,0.00039307025,0.0013245604,0.13481763,0.006481109,0.045867596],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99830794,0.0003456918,0.00016728588,0.0005543989,0.00046125945,0.00016338672],"domain_scores_gemma":[0.9973591,0.00066968193,0.00013154319,0.0005555262,0.0010684686,0.00021568958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026104082,0.0014283498,0.0012944562,0.0015950827,0.00076412264,0.0018819735,0.0025130336,0.0019359528,0.025387315],"category_scores_gemma":[0.0052867774,0.00065464416,0.001214992,0.00095015636,0.00037748006,0.0030902808,0.0011855344,0.0014836059,0.020757211],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003956104,0.0013341096,0.003154346,0.0012224172,0.00027761486,0.0010319365,0.0006092626,0.0073128417,0.13201605,0.003300585,0.33102044,0.51476425],"study_design_scores_gemma":[0.001480849,0.003305182,0.021644402,0.00014159038,0.0004206031,0.003177905,0.0007110099,0.4510662,0.29521692,0.007427775,0.21493852,0.0004691938],"about_ca_topic_score_codex":0.0050188587,"about_ca_topic_score_gemma":0.0047172024,"teacher_disagreement_score":0.025387315,"about_ca_system_score_codex":0.00086965493,"about_ca_system_score_gemma":0.0012463839,"threshold_uncertainty_score":0.08492905},"labels":[],"label_agreement":null},{"id":"W2296352097","doi":"10.1109/icmla.2015.142","title":"Active Information Retrieval for Linking Twitter Posts with Political Debates","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada; International Development Research Centre","keywords":"Computer science; Information retrieval; Microblogging; Oracle; Set (abstract data type); Process (computing); Social media; Task (project management); Selection (genetic algorithm); Key (lock); Event (particle physics); World Wide Web; Artificial intelligence","score_opus":0.04041231629114066,"score_gpt":0.2738708666052466,"score_spread":0.23345855031410595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296352097","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.070887804,0.0016622752,0.9195593,0.00072064326,0.00011519037,0.0004886462,0.00063257717,0.0027159057,0.0032176538],"genre_scores_gemma":[0.66932464,0.00071226776,0.32198378,0.00027148443,0.00049066974,0.0005717274,0.001947277,0.00016998185,0.0045281462],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99776876,0.0009658304,0.00015536376,0.0004950786,0.00045266782,0.00016233607],"domain_scores_gemma":[0.9935231,0.00447257,0.0006029847,0.00063838926,0.000607333,0.00015553716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039095073,0.0011658774,0.0011705841,0.0061597344,0.0012656497,0.002513244,0.002591078,0.0024497397,0.0030726723],"category_scores_gemma":[0.015093149,0.0006775651,0.0013349627,0.0036887375,0.000999406,0.005713113,0.0017410475,0.0013518906,0.0016418283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001928696,0.0012103658,0.006538295,0.0010312992,0.00040424272,0.0005149287,0.0018951544,0.11294089,0.035701346,0.03490449,0.0123729715,0.79055727],"study_design_scores_gemma":[0.00007961261,0.00017047,0.0016193268,0.000035415484,0.0001035675,0.00020141921,0.00019271903,0.9661401,0.011159393,0.016312826,0.003933156,0.000052005737],"about_ca_topic_score_codex":0.003424886,"about_ca_topic_score_gemma":0.0037160746,"teacher_disagreement_score":0.0061597344,"about_ca_system_score_codex":0.0012045056,"about_ca_system_score_gemma":0.00095723424,"threshold_uncertainty_score":0.020675719},"labels":[],"label_agreement":null},{"id":"W2296383052","doi":"10.3115/v1/n15-1065","title":"Expanding Paraphrase Lexicons by Exploiting Lexical Variants","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Paraphrase; Computer science; Natural language processing; Artificial intelligence; Linguistics","score_opus":0.09131101412097056,"score_gpt":0.3026592437913347,"score_spread":0.21134822967036412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296383052","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23220311,0.0013923945,0.7453859,0.00035904645,0.00007262333,0.000568575,0.0019604047,0.010228017,0.007829939],"genre_scores_gemma":[0.5009867,0.00071753777,0.48435208,0.00020271602,0.00009766203,0.00023579046,0.008925318,0.0012883529,0.0031938718],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990339,0.00023813022,0.00013057646,0.00030736244,0.00021649398,0.00007351766],"domain_scores_gemma":[0.9974482,0.0012291007,0.00020161689,0.0005700639,0.00048386602,0.000067172725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007688935,0.001115342,0.000889489,0.0038781345,0.00054367614,0.0012116601,0.0011050438,0.00063862297,0.0031056774],"category_scores_gemma":[0.0053486363,0.00077401823,0.0010666933,0.0025162043,0.00039051115,0.0026741072,0.0017873275,0.0008424886,0.0026669505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022497811,0.00028545357,0.0066150827,0.00072295417,0.00015967218,0.0012763272,0.0010062753,0.011132208,0.14885221,0.0055262838,0.008254082,0.81594443],"study_design_scores_gemma":[0.00020420608,0.00067994464,0.017409375,0.0002237089,0.00064062356,0.007701869,0.0014335616,0.6805932,0.19730069,0.03422236,0.0594037,0.00018677198],"about_ca_topic_score_codex":0.0015027504,"about_ca_topic_score_gemma":0.0035039813,"teacher_disagreement_score":0.0038781345,"about_ca_system_score_codex":0.00034499838,"about_ca_system_score_gemma":0.00074532867,"threshold_uncertainty_score":0.010389507},"labels":[],"label_agreement":null},{"id":"W2296397244","doi":"","title":"UCD IIRG at TAC KBP 2013","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Salient; Computer science; Knowledge base; Linkage (software); Population; Base (topology); Artificial intelligence; Mathematics; Genetics; Gene; Biology","score_opus":0.00899483821476027,"score_gpt":0.22776748656930118,"score_spread":0.2187726483545409,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296397244","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03598129,0.0026853348,0.27435306,0.007838359,0.0041261273,0.0027048942,0.14616412,0.28565934,0.24048755],"genre_scores_gemma":[0.10409745,0.00093230983,0.32314736,0.0018962087,0.0006212492,0.0017687484,0.40392166,0.026764605,0.13685046],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9942163,0.0013663225,0.00025920456,0.0010523135,0.002681767,0.00042400532],"domain_scores_gemma":[0.9912157,0.0016233315,0.00026277135,0.002209207,0.003885945,0.0008029623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009597837,0.0016458177,0.0017188682,0.0034750851,0.0029353062,0.008379015,0.0036582767,0.0021672351,0.07251372],"category_scores_gemma":[0.020718757,0.0011574213,0.0009064168,0.0034229497,0.0011257344,0.007485599,0.0036947748,0.0035389941,0.05738105],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007415574,0.0004001083,0.001038286,0.0002953637,0.00006492104,0.0006539925,0.0005446454,0.0027157138,0.0043129297,0.00986067,0.86648256,0.1128893],"study_design_scores_gemma":[0.0008098068,0.00026375422,0.0026308638,0.00016963559,0.000110276866,0.0005001525,0.00047803944,0.06681541,0.013332175,0.010788274,0.90395826,0.00014342257],"about_ca_topic_score_codex":0.06670112,"about_ca_topic_score_gemma":0.052571896,"teacher_disagreement_score":0.07251372,"about_ca_system_score_codex":0.004502276,"about_ca_system_score_gemma":0.005138285,"threshold_uncertainty_score":0.24258262},"labels":[],"label_agreement":null},{"id":"W2296508364","doi":"","title":"Optimizing Question-Answering Systems Using Genetic Algorithms","year":2015,"lang":"en","type":"article","venue":"The Florida AI Research Society","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Defence Research and Development Canada; Lakehead University; Université Laval","funders":"","keywords":"Adaptability; Sequence (biology); Computer science; Question answering; Genetic algorithm; Algorithm; Space (punctuation); Artificial intelligence; Machine learning; Theoretical computer science","score_opus":0.17650172490538663,"score_gpt":0.3937401107580525,"score_spread":0.21723838585266586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296508364","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18123429,0.0007681852,0.8094946,0.0008114983,0.000052184754,0.0002730934,0.00009558436,0.0019119955,0.0053586173],"genre_scores_gemma":[0.70770466,0.00040383555,0.28879577,0.00030528227,0.000036467256,0.00031676883,0.00028066087,0.00017613429,0.001980468],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988086,0.00060501,0.00005791552,0.00024352994,0.00017210543,0.00011290168],"domain_scores_gemma":[0.9975497,0.0018815801,0.0001435948,0.00012380558,0.00025189106,0.00004938953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020212615,0.00092996075,0.0008771607,0.0009094814,0.0005924843,0.0012930145,0.0011609315,0.001636148,0.0014949804],"category_scores_gemma":[0.0068687433,0.00047430958,0.00056883984,0.0006612508,0.0009059009,0.0013795964,0.0008002151,0.00087383756,0.00037129037],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006058876,0.000087145614,0.0012873302,0.000058557984,0.000054343276,0.000054615317,0.000119758275,0.95492125,0.0032676258,0.005047838,0.0005420133,0.03449885],"study_design_scores_gemma":[0.000016225687,0.000028541079,0.00012362306,0.0000051184948,0.000014578524,0.000010713556,0.000034157343,0.9940853,0.00096935837,0.004319511,0.00038692585,0.000005926792],"about_ca_topic_score_codex":0.0062428354,"about_ca_topic_score_gemma":0.0059084585,"teacher_disagreement_score":0.0062428354,"about_ca_system_score_codex":0.0013175174,"about_ca_system_score_gemma":0.0012167571,"threshold_uncertainty_score":0.012413025},"labels":[],"label_agreement":null},{"id":"W2296555785","doi":"10.1007/978-3-319-25789-1_10","title":"Semi-extractive Multi-document Summarization via Submodular Functions","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Submodular set function; Automatic summarization; Computer science; Monotone polygon; Knapsack problem; Scalability; Relevance (law); Greedy algorithm; Theoretical computer science; Data mining; Mathematical optimization; Artificial intelligence; Algorithm; Mathematics; Database","score_opus":0.030237387980650392,"score_gpt":0.2617135266681185,"score_spread":0.23147613868746808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296555785","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004005824,0.0010478013,0.9892675,0.0002057014,0.00011101416,0.000091349975,0.00075019826,0.0033545617,0.0011660904],"genre_scores_gemma":[0.08654568,0.0014063722,0.89012843,0.00028883753,0.0005161165,0.00043770674,0.008033839,0.0009436954,0.011699306],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983022,0.0004410608,0.00019130608,0.00045340488,0.00046440566,0.00014761518],"domain_scores_gemma":[0.997071,0.0015165054,0.0002033268,0.00044958372,0.0006697979,0.00008971794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015633323,0.0023298042,0.002516372,0.0025377125,0.0007305469,0.0023587067,0.0017650619,0.0014466178,0.00621643],"category_scores_gemma":[0.0046631503,0.00086784334,0.0018552319,0.0029878633,0.0005770582,0.003593053,0.0026512824,0.002000382,0.008002696],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035278435,0.0001495032,0.00024409489,0.00064096926,0.00014786984,0.00021751116,0.00020274926,0.038766533,0.024568217,0.008271413,0.021508401,0.90492994],"study_design_scores_gemma":[0.000058330552,0.00027425037,0.00046031273,0.00008442339,0.0001424469,0.00030623475,0.00022062185,0.92345226,0.023325086,0.035259746,0.01635261,0.00006370644],"about_ca_topic_score_codex":0.0012722113,"about_ca_topic_score_gemma":0.0027278531,"teacher_disagreement_score":0.00621643,"about_ca_system_score_codex":0.0005941868,"about_ca_system_score_gemma":0.00129851,"threshold_uncertainty_score":0.02079606},"labels":[],"label_agreement":null},{"id":"W2296709160","doi":"","title":"GUCAS at TREC 2011 Microblog Track.","year":2011,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Microblogging; Social media; Context (archaeology); Relevance (law); Track (disk drive); Probabilistic logic; Information retrieval; Baseline (sea); Natural language processing; Query expansion; Artificial intelligence; Language model; World Wide Web","score_opus":0.09059576231893875,"score_gpt":0.25523591573740934,"score_spread":0.1646401534184706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296709160","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10564944,0.014168375,0.021209337,0.014621,0.017034084,0.006981912,0.665282,0.07146175,0.08359205],"genre_scores_gemma":[0.081927255,0.0026904528,0.041468475,0.0021497733,0.0020000604,0.0033739563,0.7520476,0.0023321419,0.11201039],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952023,0.0016081705,0.00024301517,0.00066723936,0.0017957914,0.00048356078],"domain_scores_gemma":[0.9892335,0.0024838143,0.00036847306,0.0015114022,0.0047454885,0.001657298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073954742,0.003923351,0.002864947,0.006008905,0.0039439905,0.0038228033,0.002702709,0.002966729,0.03237714],"category_scores_gemma":[0.014871695,0.0007151858,0.0011093096,0.0035324304,0.00065842504,0.006441104,0.002290321,0.0030527695,0.01968764],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034605231,0.00069485087,0.0006904686,0.0005582514,0.000089696536,0.00008174958,0.00006531863,0.0011132892,0.0021312684,0.00053423864,0.9588061,0.034888797],"study_design_scores_gemma":[0.0019282438,0.0015311657,0.021685459,0.00030478975,0.00041038904,0.0003830646,0.0010437481,0.096178755,0.03209002,0.0033285657,0.84071696,0.00039895065],"about_ca_topic_score_codex":0.17701058,"about_ca_topic_score_gemma":0.24732497,"teacher_disagreement_score":0.17701058,"about_ca_system_score_codex":0.006484662,"about_ca_system_score_gemma":0.0055450397,"threshold_uncertainty_score":0.35196054},"labels":[],"label_agreement":null},{"id":"W2302733396","doi":"10.15837/ijccc.2016.3.700","title":"Efficient Opinion Summarization on Comments with Online-LDA","year":2016,"lang":"en","type":"article","venue":"International Journal of Computers Communications & Control","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Computer science; Variety (cybernetics); Information retrieval; Set (abstract data type); Data science; World Wide Web; Artificial intelligence","score_opus":0.02421673091146074,"score_gpt":0.29234242506061325,"score_spread":0.2681256941491525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2302733396","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04350726,0.0019557306,0.9430548,0.0007941471,0.0004084176,0.00034223808,0.0017788307,0.005476815,0.002681804],"genre_scores_gemma":[0.53352946,0.0013089932,0.44319627,0.000404942,0.0018480725,0.0005632316,0.009982767,0.0003933554,0.008772918],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982377,0.0005324383,0.00014640643,0.0004286393,0.00046158783,0.00019326128],"domain_scores_gemma":[0.99710006,0.00085169595,0.00028070487,0.0003272999,0.001344221,0.00009593113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017365218,0.0016981475,0.001506305,0.0036268996,0.0008548793,0.0015703173,0.0010748556,0.0008702728,0.0015251653],"category_scores_gemma":[0.005430244,0.0003252209,0.0012905132,0.0022606268,0.0003327765,0.0020567237,0.0011503717,0.0012236568,0.0026300612],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056263787,0.00039100047,0.0041853487,0.00040263214,0.0002676428,0.00014927925,0.0006213453,0.017885203,0.03019988,0.0017961811,0.024482753,0.91905606],"study_design_scores_gemma":[0.00007644425,0.0001928092,0.0037336764,0.00003159879,0.00020395147,0.00011619053,0.0005442114,0.95806265,0.02094248,0.0056447308,0.010388271,0.000062945335],"about_ca_topic_score_codex":0.0056663794,"about_ca_topic_score_gemma":0.009730396,"teacher_disagreement_score":0.0056663794,"about_ca_system_score_codex":0.0006388729,"about_ca_system_score_gemma":0.001017266,"threshold_uncertainty_score":0.011266768},"labels":[],"label_agreement":null},{"id":"W2304545146","doi":"10.18653/v1/p16-1056","title":"Generating Factoid Questions With Recurrent Neural Networks: The 30M Factoid Question-Answer Corpus","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Canadian Institute for Advanced Research","keywords":"Computer science; Natural language processing; Artificial intelligence; Question answering; Sentence; Similarity (geometry); Machine translation; Knowledge base; Baseline (sea); Architecture; Artificial neural network; Information retrieval","score_opus":0.027350154643534606,"score_gpt":0.26779360357769205,"score_spread":0.24044344893415745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2304545146","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61682165,0.009988075,0.12471778,0.007814707,0.0021423155,0.0012946853,0.18868203,0.018444933,0.030093867],"genre_scores_gemma":[0.5213131,0.00086627604,0.10951143,0.0006299651,0.00031125767,0.00088497024,0.35294566,0.0009521234,0.012585181],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982311,0.00086561596,0.00009689498,0.00048432755,0.00023977598,0.00008228234],"domain_scores_gemma":[0.994624,0.0036668733,0.00014915531,0.0006941855,0.00066687487,0.00019891042],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016825501,0.00133309,0.0008562421,0.0016205219,0.0012928335,0.0013839535,0.0016169585,0.0023743561,0.01181944],"category_scores_gemma":[0.0134358,0.0005159985,0.0007973953,0.0012746886,0.0008135154,0.0030507138,0.0025022526,0.0019984657,0.006595902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002246863,0.0012127704,0.011417926,0.0038414083,0.0003652044,0.0028964188,0.005277475,0.02179639,0.022903465,0.019963136,0.47508225,0.43299675],"study_design_scores_gemma":[0.0010462407,0.00074609224,0.031344444,0.0006877183,0.00029305532,0.0026635793,0.0053914106,0.41619167,0.03513747,0.04716838,0.45905492,0.00027504837],"about_ca_topic_score_codex":0.008836431,"about_ca_topic_score_gemma":0.014323425,"teacher_disagreement_score":0.01181944,"about_ca_system_score_codex":0.0011982352,"about_ca_system_score_gemma":0.0012266956,"threshold_uncertainty_score":0.039539993},"labels":[],"label_agreement":null},{"id":"W2326533993","doi":"10.48550/arxiv.1603.09025","title":"Recurrent Batch Normalization","year":2016,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Normalization (sociology); Computer science; Recurrent neural network; Batch processing; Transformation (genetics); Artificial intelligence; Generalization; Sequence (biology); Convergence (economics); Artificial neural network; Machine learning; Mathematics","score_opus":0.08285591806591706,"score_gpt":0.18734298762263468,"score_spread":0.10448706955671762,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2326533993","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010323771,0.00028420726,0.9814873,0.0003422027,0.00026084055,0.00007336998,0.0003617724,0.0038232089,0.0030434167],"genre_scores_gemma":[0.47128573,0.0005682298,0.50504845,0.0009423464,0.00045855352,0.0005787068,0.002098045,0.0014864184,0.017533476],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99880636,0.00026662482,0.00007774008,0.00047701164,0.000264503,0.00010778535],"domain_scores_gemma":[0.9980829,0.00047443976,0.00016780768,0.00076317764,0.00043136766,0.00008029952],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016022346,0.0013740943,0.0012026414,0.0005274754,0.00051966833,0.0011954255,0.0035260597,0.0012483383,0.008423565],"category_scores_gemma":[0.007026495,0.00057454733,0.0011392056,0.0008880727,0.0009773332,0.00349954,0.0016969208,0.0026223375,0.0039502387],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005448519,0.00035402525,0.0022054594,0.0003084814,0.00024503958,0.00030614555,0.00037906066,0.24766362,0.065844424,0.07471293,0.025022693,0.58241326],"study_design_scores_gemma":[0.000022041955,0.0000927548,0.00044329403,0.000017627714,0.000048356233,0.00010222448,0.000025135641,0.94006824,0.019431286,0.032940667,0.006775536,0.00003281465],"about_ca_topic_score_codex":0.004208412,"about_ca_topic_score_gemma":0.0067300717,"teacher_disagreement_score":0.008423565,"about_ca_system_score_codex":0.0011467001,"about_ca_system_score_gemma":0.0017671881,"threshold_uncertainty_score":0.028179646},"labels":[],"label_agreement":null},{"id":"W2331470947","doi":"10.3758/bf03192768","title":"Word frequency effects in high-dimensional co-occurrence models: A new approach","year":2006,"lang":"en","type":"article","venue":"Behavior Research Methods","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Word (group theory); Word lists by frequency; Computer science; Distributional semantics; Lexical decision task; Natural language processing; Orthographic projection; Hyperspace; Space (punctuation); Semantics (computer science); Artificial intelligence; Linguistics; Semantic similarity; Cognition; Psychology","score_opus":0.30327505621301126,"score_gpt":0.5142418890998628,"score_spread":0.21096683288685153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2331470947","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014414026,0.0005556943,0.98349285,0.00038761727,0.000059274735,0.000076747376,0.00017450147,0.00022502926,0.00061430375],"genre_scores_gemma":[0.44292068,0.0017708899,0.54785246,0.0004375433,0.00084627926,0.0013749358,0.0008602508,0.0004003486,0.0035366644],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9770791,0.01650332,0.0011076422,0.0032293936,0.0016705523,0.00040993927],"domain_scores_gemma":[0.80618966,0.17934597,0.0026437226,0.008640738,0.002298601,0.0008812522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028588837,0.0020882634,0.0040076296,0.0056910343,0.0019371593,0.0072782687,0.005906907,0.0036283557,0.0049000666],"category_scores_gemma":[0.08261272,0.0018755605,0.0048767873,0.007410039,0.0037023406,0.008226249,0.0043970305,0.006466056,0.0010546227],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012607138,0.0017600077,0.034527764,0.0011385815,0.0049279407,0.0008167482,0.0065555377,0.1699386,0.0051739966,0.47315365,0.0042324606,0.29651415],"study_design_scores_gemma":[0.00013419136,0.00023716052,0.0046564993,0.00009429385,0.00092201645,0.00025591545,0.000444775,0.75058776,0.0006427091,0.23865995,0.0032266274,0.0001380534],"about_ca_topic_score_codex":0.005456197,"about_ca_topic_score_gemma":0.005827234,"teacher_disagreement_score":0.028588837,"about_ca_system_score_codex":0.0015282967,"about_ca_system_score_gemma":0.0022318328,"threshold_uncertainty_score":0.15119398},"labels":[],"label_agreement":null},{"id":"W2341010248","doi":"10.1177/0163278715605358","title":"Using Automated Scoring to Evaluate Written Responses in English and French on a High-Stakes Clinical Competency Examination","year":2015,"lang":"en","type":"article","venue":"Evaluation & the Health Professions","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Medical Council of Canada; University of Alberta","funders":"","keywords":"Reliability (semiconductor); Computer science; Stage (stratigraphy); Artificial intelligence; Natural language processing; Machine learning; Medical education; Medicine","score_opus":0.4259430117257569,"score_gpt":0.5111785923759283,"score_spread":0.08523558065017145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2341010248","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6752609,0.00026206634,0.31115165,0.0003782534,0.000084145344,0.0013281782,0.0011179007,0.0029546642,0.0074623763],"genre_scores_gemma":[0.8023966,0.00011047092,0.19365028,0.00009104976,0.00005532152,0.0006816608,0.0012379544,0.000081060156,0.0016955002],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9780075,0.015243071,0.0011838032,0.001559166,0.0034731831,0.0005332525],"domain_scores_gemma":[0.94899255,0.029779624,0.00484361,0.0025517584,0.013023042,0.0008094408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01581016,0.0009760304,0.000615637,0.0041752513,0.0006104189,0.0017592685,0.00085164106,0.0007321174,0.0014886765],"category_scores_gemma":[0.04321626,0.00021185991,0.00052771426,0.0017047095,0.0004591133,0.00084376644,0.0014613988,0.00059934845,0.00095676736],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060364464,0.0006139239,0.20006397,0.0003416812,0.0002860421,0.00023050136,0.0035822382,0.018117178,0.02887038,0.0015056113,0.004552653,0.74123216],"study_design_scores_gemma":[0.00023356895,0.0016788689,0.38499275,0.00018933832,0.00023782569,0.0009300447,0.0038607027,0.5331036,0.05838478,0.004374868,0.011615702,0.00039784028],"about_ca_topic_score_codex":0.010518738,"about_ca_topic_score_gemma":0.021808386,"teacher_disagreement_score":0.01581016,"about_ca_system_score_codex":0.0011125676,"about_ca_system_score_gemma":0.0019218543,"threshold_uncertainty_score":0.0836131},"labels":[],"label_agreement":null},{"id":"W2342771006","doi":"","title":"Measuring Semantic Similarity using a Multi-Tree Model","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Semantic similarity; Similarity (geometry); Tree (set theory); Artificial intelligence; Natural language processing; Mathematics","score_opus":0.3337684261450095,"score_gpt":0.28177573519248744,"score_spread":0.05199269095252207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2342771006","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038231272,0.00064483506,0.95499116,0.00033849603,0.00004811705,0.00017211343,0.00038121143,0.0004385959,0.004754169],"genre_scores_gemma":[0.4774237,0.0008098835,0.51780206,0.00014949952,0.00009740058,0.00034191948,0.0011396696,0.00015705265,0.0020788247],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99459416,0.0019394964,0.0003735364,0.00090196595,0.0019624298,0.0002284077],"domain_scores_gemma":[0.9919837,0.004815877,0.00078385556,0.00086229626,0.0013334595,0.00022081278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033466234,0.00066366413,0.0013657615,0.007481144,0.0011911755,0.0032653133,0.0021577012,0.002689558,0.0022567788],"category_scores_gemma":[0.018998742,0.00048154435,0.0016808145,0.007879677,0.0010361193,0.011327994,0.0020257575,0.001496384,0.0010845885],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037021356,0.00047519503,0.014851141,0.0005728082,0.00063388736,0.00047667834,0.0018285359,0.39733827,0.0099985,0.2636784,0.005880696,0.30389562],"study_design_scores_gemma":[0.000012878981,0.00007934121,0.001775144,0.00003079062,0.000051782798,0.00023357202,0.00013447728,0.88685,0.00087794545,0.10721972,0.0027023877,0.000032006064],"about_ca_topic_score_codex":0.0052830074,"about_ca_topic_score_gemma":0.005010781,"teacher_disagreement_score":0.007481144,"about_ca_system_score_codex":0.0019171314,"about_ca_system_score_gemma":0.0010979458,"threshold_uncertainty_score":0.017698824},"labels":[],"label_agreement":null},{"id":"W2356597680","doi":"","title":"Named entity recognition in Chinese medical records based on cascaded conditional random field","year":2014,"lang":"en","type":"article","venue":"Journal of Jilin University","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Conditional random field; CRFS; Feature (linguistics); Named-entity recognition; Context (archaeology); Computer science; Sentence; Word (group theory); Pattern recognition (psychology); Artificial intelligence; Layer (electronics); Natural language processing; Field (mathematics); Speech recognition; Mathematics; Engineering; Linguistics","score_opus":0.010216477530157586,"score_gpt":0.22386010737798961,"score_spread":0.21364362984783203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2356597680","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07221107,0.0010018894,0.9157824,0.00038712198,0.00014854466,0.0003595511,0.0027168326,0.0062259706,0.0011666607],"genre_scores_gemma":[0.56641006,0.00086218375,0.42092052,0.00015697909,0.00020227753,0.00031436147,0.008249048,0.00015170757,0.0027329028],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985751,0.00027786617,0.00015500901,0.0005780193,0.00029635368,0.00011763395],"domain_scores_gemma":[0.9974234,0.0014544143,0.0002636988,0.0003368516,0.0004469013,0.00007471545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021234187,0.0008858155,0.00093561533,0.0026490705,0.0006144793,0.00054740516,0.0014524438,0.0007763577,0.0017926294],"category_scores_gemma":[0.004151581,0.00037490748,0.0015047223,0.0022242556,0.00030186912,0.002034542,0.00069975614,0.0007877462,0.00075487595],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009063643,0.00030868975,0.026192382,0.00072414987,0.00035295478,0.0020173956,0.0006694928,0.08506537,0.037991595,0.0058196345,0.021708172,0.81824386],"study_design_scores_gemma":[0.000047694513,0.00014516522,0.014638884,0.000034723158,0.00018435923,0.00089591503,0.000082934916,0.95687395,0.018172592,0.0039670016,0.0048642885,0.000092418646],"about_ca_topic_score_codex":0.01606014,"about_ca_topic_score_gemma":0.014482283,"teacher_disagreement_score":0.01606014,"about_ca_system_score_codex":0.0006750186,"about_ca_system_score_gemma":0.0014899289,"threshold_uncertainty_score":0.031933308},"labels":[],"label_agreement":null},{"id":"W2368478439","doi":"","title":"Text Summarization Based on the Sentence Features and Semantic Distance","year":2009,"lang":"en","type":"article","venue":"Microcomputer applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Salient; Natural language processing; Feature (linguistics); Artificial intelligence; Multi-document summarization; Semantic feature; Information retrieval; Linguistics","score_opus":0.008368466798750765,"score_gpt":0.22061219190730905,"score_spread":0.21224372510855827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2368478439","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03537788,0.0010651215,0.959129,0.00018531097,0.000128157,0.00012419128,0.00038656726,0.0019620387,0.0016417345],"genre_scores_gemma":[0.33758014,0.0008929749,0.655396,0.000081652805,0.00044951937,0.00025504656,0.002243021,0.0003115989,0.0027900713],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944156,0.00013036194,0.0000636885,0.00012698602,0.00020361865,0.00003368641],"domain_scores_gemma":[0.9989869,0.00031147277,0.00015201105,0.00007973417,0.00043405508,0.000035813166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005263406,0.0008660293,0.00086145825,0.0029228565,0.00039114212,0.0007463561,0.0005305876,0.0003998183,0.0014588437],"category_scores_gemma":[0.0024897903,0.00021370572,0.0007368849,0.0017279505,0.0002657687,0.0019028218,0.0004806987,0.00053825695,0.00084432337],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029606614,0.00008581337,0.0015334543,0.0005101864,0.00015841493,0.00016404914,0.00029980714,0.008918321,0.091991425,0.0054840613,0.005223235,0.88533515],"study_design_scores_gemma":[0.0001335972,0.0013339714,0.021124765,0.00011483589,0.0007179462,0.00093620154,0.00063457096,0.76906115,0.14355195,0.028539006,0.033607036,0.00024498813],"about_ca_topic_score_codex":0.00072290323,"about_ca_topic_score_gemma":0.00097504794,"teacher_disagreement_score":0.0029228565,"about_ca_system_score_codex":0.00026777678,"about_ca_system_score_gemma":0.00038187118,"threshold_uncertainty_score":0.0048803687},"labels":[],"label_agreement":null},{"id":"W2394737088","doi":"","title":"Global and Local Models for Multi-Document Summarization.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Multi-document summarization; Information retrieval","score_opus":0.029297575181961445,"score_gpt":0.2665994952014523,"score_spread":0.23730192001949088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2394737088","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010794562,0.007272781,0.97538096,0.0013516526,0.00027763977,0.0001111812,0.0010938189,0.0010473208,0.0026700532],"genre_scores_gemma":[0.550876,0.0061023007,0.41081586,0.000804325,0.002447924,0.00085218257,0.0076142484,0.00065172574,0.019835511],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971566,0.0015342709,0.00018733951,0.00056232495,0.00037944014,0.00017997477],"domain_scores_gemma":[0.9909294,0.006364495,0.0005466271,0.0009831085,0.0009537934,0.00022267434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048680804,0.0009566484,0.0015796279,0.002814502,0.00065436977,0.002472937,0.0025631862,0.0017286539,0.004072272],"category_scores_gemma":[0.015798263,0.000532735,0.0014081588,0.0026507478,0.00081192283,0.005296916,0.0013672158,0.0021079008,0.002220486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067823625,0.00026282895,0.005626301,0.0014325613,0.0009127942,0.0003237612,0.0017403257,0.2115227,0.005177858,0.16836822,0.049995657,0.5539589],"study_design_scores_gemma":[0.000038689297,0.00011389063,0.0016468813,0.00008183809,0.00021325352,0.0001583844,0.00019872333,0.8136744,0.0011166367,0.17085019,0.011855104,0.000052031417],"about_ca_topic_score_codex":0.004162173,"about_ca_topic_score_gemma":0.00807512,"teacher_disagreement_score":0.0048680804,"about_ca_system_score_codex":0.0010372457,"about_ca_system_score_gemma":0.00081481686,"threshold_uncertainty_score":0.025745153},"labels":[],"label_agreement":null},{"id":"W2394859052","doi":"","title":"HITS' Monolingual and Cross-lingual Entity Linking System at TAC 2012: A Joint Approach.","year":2012,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Cluster analysis; Natural language processing; Joint (building); Entity linking; Knowledge base","score_opus":0.0180952552988818,"score_gpt":0.2577802106315116,"score_spread":0.2396849553326298,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2394859052","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11153712,0.0021948528,0.4655939,0.002078594,0.0012891575,0.0012151279,0.02534245,0.3506535,0.040095303],"genre_scores_gemma":[0.31719917,0.0005426901,0.4958932,0.0010498378,0.0003633667,0.0007942648,0.13129577,0.00877202,0.04408969],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99613285,0.0008968799,0.0003861634,0.0011025316,0.0010682291,0.00041335396],"domain_scores_gemma":[0.993338,0.0010137052,0.0002999795,0.002203685,0.0025029066,0.00064174016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041089826,0.0018251167,0.0014167724,0.0038717717,0.0020703096,0.0031468002,0.0031196168,0.0027471192,0.009483169],"category_scores_gemma":[0.0077219456,0.0011155184,0.0015536974,0.0031755387,0.00060620427,0.009069387,0.0057788487,0.0020989159,0.014800927],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011408095,0.0009771397,0.010211489,0.0010254412,0.0010313055,0.001856111,0.0017483213,0.012109593,0.044045426,0.006150917,0.42221898,0.49748442],"study_design_scores_gemma":[0.00039027678,0.0011157384,0.019327834,0.00020282544,0.0010886992,0.0032611059,0.003060621,0.43114248,0.14061603,0.01781726,0.38120955,0.0007675839],"about_ca_topic_score_codex":0.015251458,"about_ca_topic_score_gemma":0.02399156,"teacher_disagreement_score":0.015251458,"about_ca_system_score_codex":0.0011540239,"about_ca_system_score_gemma":0.003655042,"threshold_uncertainty_score":0.031724334},"labels":[],"label_agreement":null},{"id":"W2395375875","doi":"","title":"Coming to Terms: A Discourse Epistemetrics Study of Article Abstracts from the Web of Science.","year":2015,"lang":"en","type":"article","venue":"ISSI","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Discipline; Metadiscourse; Adverbial; Framing (construction); Computer science; Set (abstract data type); Data science; Psychology; Linguistics; Artificial intelligence; Sociology; Social science","score_opus":0.07798617127179955,"score_gpt":0.3311430903026429,"score_spread":0.25315691903084336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395375875","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91630274,0.009692752,0.035121012,0.002914486,0.0001841095,0.00036831928,0.0048473724,0.00015869536,0.0304106],"genre_scores_gemma":[0.9697337,0.0027416523,0.018591462,0.00021860765,0.00018722306,0.00051292195,0.00439206,0.00011366499,0.0035085836],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99258876,0.0035976064,0.00082020334,0.0008035273,0.001979869,0.00021001878],"domain_scores_gemma":[0.9399558,0.04761338,0.0054510506,0.001857843,0.00390595,0.0012159795],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.008457203,0.00048342586,0.00055361306,0.024723366,0.003181107,0.0061171195,0.0007009626,0.0010151191,0.0024486517],"category_scores_gemma":[0.06728952,0.00024715794,0.0005486327,0.032541268,0.0028790748,0.008058172,0.0041343668,0.0011831541,0.0006735441],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008119792,0.00035413844,0.1639239,0.005369796,0.00030097287,0.0018958439,0.43078166,0.0020969426,0.015861068,0.087023936,0.010886662,0.28069308],"study_design_scores_gemma":[0.00009543185,0.00028001593,0.26591238,0.0021582218,0.00038694136,0.0024993217,0.32391647,0.01553336,0.009622871,0.08149334,0.297865,0.00023662312],"about_ca_topic_score_codex":0.0045256037,"about_ca_topic_score_gemma":0.006617715,"teacher_disagreement_score":0.9915428,"about_ca_system_score_codex":0.0029207726,"about_ca_system_score_gemma":0.0024209714,"threshold_uncertainty_score":0.04472649},"labels":[],"label_agreement":null},{"id":"W2395546452","doi":"","title":"Balanced Coverage of Aspects for Text Summarization.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Task (project management); Benchmark (surveying); Principle of maximum entropy; Artificial intelligence; Classifier (UML); Information retrieval; Data mining; Natural language processing","score_opus":0.015219384811209108,"score_gpt":0.23375145413788723,"score_spread":0.21853206932667812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395546452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023566447,0.0017741835,0.96810305,0.00055053073,0.000074474054,0.00022502089,0.0012864552,0.0022576486,0.0021621843],"genre_scores_gemma":[0.5501298,0.0010757311,0.43372762,0.00032904063,0.00043969738,0.00091109076,0.008928874,0.00043615792,0.004021968],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977671,0.0009345603,0.00017707232,0.00052590435,0.000497895,0.00009747736],"domain_scores_gemma":[0.9939673,0.0037979865,0.00082634646,0.0005671732,0.00069466,0.00014644268],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024883377,0.0012891729,0.001315597,0.0026912903,0.000562932,0.0015613032,0.0014569847,0.0014389134,0.0014818981],"category_scores_gemma":[0.012959481,0.00044096468,0.0011821954,0.0023147287,0.00046593425,0.003247806,0.0009933203,0.0010877304,0.0009620219],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059857493,0.00023320537,0.009278167,0.0011828906,0.00039215278,0.0004802239,0.0012290688,0.30169266,0.015218753,0.023820972,0.023812488,0.62206084],"study_design_scores_gemma":[0.000028223505,0.00016134274,0.0020091687,0.000058259637,0.00012118451,0.00020253789,0.000104061386,0.9569897,0.00406592,0.028542386,0.0076876436,0.000029611256],"about_ca_topic_score_codex":0.002482801,"about_ca_topic_score_gemma":0.004319433,"teacher_disagreement_score":0.0026912903,"about_ca_system_score_codex":0.0010059936,"about_ca_system_score_gemma":0.0008258534,"threshold_uncertainty_score":0.0131598115},"labels":[],"label_agreement":null},{"id":"W2395937579","doi":"","title":"QA System Metis Based on Semantic Graph Matching at NTCIR 6.","year":2007,"lang":"en","type":"article","venue":"NTCIR","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Metis; Artificial intelligence; Natural language processing; Graph; Matching (statistics); Information retrieval; World Wide Web; Theoretical computer science; Mathematics; Statistics","score_opus":0.015687917336282126,"score_gpt":0.2348742655769794,"score_spread":0.21918634824069727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2395937579","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050563514,0.0014332449,0.6749339,0.0023277022,0.000907688,0.0020589188,0.027138967,0.17156003,0.06907608],"genre_scores_gemma":[0.38995475,0.0004046786,0.50161636,0.00046357536,0.00029970947,0.0007946591,0.07597366,0.0056974897,0.024795199],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969119,0.001267541,0.00023786655,0.0006497231,0.0006783804,0.00025455988],"domain_scores_gemma":[0.99591583,0.001021349,0.0001740735,0.000876023,0.001811406,0.00020144235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044835526,0.00085872423,0.0013366676,0.004336286,0.0017118509,0.0032294025,0.0015254036,0.00149607,0.02697256],"category_scores_gemma":[0.009658419,0.0006076961,0.0011750063,0.0021246143,0.0004650592,0.005165939,0.0018685057,0.0012859171,0.015818607],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024384363,0.0008324897,0.009152831,0.0017365367,0.00049786974,0.0005669512,0.0013667198,0.016536666,0.04964442,0.050060738,0.32638398,0.5407824],"study_design_scores_gemma":[0.00056219957,0.0007422489,0.009726512,0.00019489577,0.00050061854,0.0006418305,0.0010874477,0.5764121,0.09977302,0.041329034,0.26877138,0.00025874167],"about_ca_topic_score_codex":0.016579987,"about_ca_topic_score_gemma":0.0122614065,"teacher_disagreement_score":0.02697256,"about_ca_system_score_codex":0.0016485435,"about_ca_system_score_gemma":0.0036673045,"threshold_uncertainty_score":0.09023219},"labels":[],"label_agreement":null},{"id":"W2396290338","doi":"","title":"On alternative automated content evaluation measures.","year":2009,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Dependency (UML); Computer science; Generative grammar; Sentence; Natural language processing; Formalism (music); Metric (unit); Artificial intelligence; Generative model; Task (project management); Recall; Information retrieval; Machine learning; Data mining","score_opus":0.04496584988138627,"score_gpt":0.3015306740907009,"score_spread":0.2565648242093146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396290338","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30686998,0.0075116865,0.6105404,0.0013618814,0.000942879,0.0040806127,0.008237708,0.010886891,0.049567994],"genre_scores_gemma":[0.76026773,0.00046020237,0.22359702,0.0003510162,0.00045763157,0.0023192824,0.0074291853,0.00065291836,0.0044649905],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.90637964,0.057843715,0.007101253,0.005952044,0.021290315,0.0014329918],"domain_scores_gemma":[0.76059717,0.16352187,0.016552264,0.018890064,0.038028177,0.002410362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.042087395,0.0019828284,0.0015193565,0.011711176,0.0014072955,0.0056118,0.0023483639,0.0027574508,0.005080958],"category_scores_gemma":[0.15640037,0.00046653888,0.0009969377,0.0067573865,0.0013039571,0.0075186538,0.0029144487,0.0016465945,0.002251138],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0053374562,0.0018545801,0.054114178,0.00288311,0.0014517067,0.00021451531,0.0022900933,0.037503622,0.022988213,0.027495246,0.02451627,0.819351],"study_design_scores_gemma":[0.0010002465,0.009652399,0.13515432,0.0008695282,0.0011176885,0.00091593387,0.0037973782,0.6654668,0.07129948,0.050736636,0.059120912,0.00086871843],"about_ca_topic_score_codex":0.0018320533,"about_ca_topic_score_gemma":0.0025050426,"teacher_disagreement_score":0.042087395,"about_ca_system_score_codex":0.0021616903,"about_ca_system_score_gemma":0.0011929225,"threshold_uncertainty_score":0.22258204},"labels":[],"label_agreement":null},{"id":"W2396516112","doi":"","title":"Description of the Google update summarizer at TAC-2011.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Task (project management); Rank (graph theory); Information retrieval; Extension (predicate logic); Mathematics; Engineering; Programming language","score_opus":0.02196588977315654,"score_gpt":0.21828320259922354,"score_spread":0.19631731282606701,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396516112","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01940863,0.009247516,0.24650763,0.0033012466,0.002974188,0.0035329664,0.21758267,0.41981116,0.07763402],"genre_scores_gemma":[0.07162195,0.0024530035,0.2010622,0.0015340764,0.00078626926,0.0018507688,0.5885082,0.023854602,0.10832892],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990361,0.00017140171,0.00010277196,0.00017969702,0.0003880449,0.00012203961],"domain_scores_gemma":[0.99806446,0.00026783298,0.00007662107,0.0003148698,0.00094701734,0.00032915978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014150142,0.0014639986,0.00084532396,0.0029196495,0.00067996344,0.001851692,0.0020516678,0.0011582214,0.042857163],"category_scores_gemma":[0.0029931278,0.00069229124,0.0006507272,0.0026877054,0.00019700255,0.0029999986,0.001196933,0.0012921017,0.044082735],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006945252,0.00015124069,0.0013055374,0.0009972703,0.00010424183,0.00041881253,0.00030249794,0.0020410854,0.01029903,0.0026120746,0.8086104,0.17246331],"study_design_scores_gemma":[0.00020354251,0.00035870008,0.004354239,0.000094291885,0.00007386396,0.0008299352,0.00017398443,0.014412687,0.009552251,0.0020679426,0.96775645,0.00012211101],"about_ca_topic_score_codex":0.00944065,"about_ca_topic_score_gemma":0.013940856,"teacher_disagreement_score":0.042857163,"about_ca_system_score_codex":0.00082139287,"about_ca_system_score_gemma":0.0009273059,"threshold_uncertainty_score":0.14337152},"labels":[],"label_agreement":null},{"id":"W2396906492","doi":"","title":"An approach using Named Entities for Recognizing Textual Entailment.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Textual entailment; Logical consequence; Computer science; Natural language processing; Artificial intelligence","score_opus":0.03825402353764714,"score_gpt":0.2759150633261194,"score_spread":0.23766103978847228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396906492","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012937989,0.0023396858,0.9700047,0.0013166878,0.00036151227,0.0006090245,0.0027231302,0.0045782514,0.0051289825],"genre_scores_gemma":[0.15927768,0.0011238913,0.8230266,0.0005115444,0.00035840555,0.000760613,0.01000969,0.00040928563,0.004522253],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9951321,0.0019983926,0.0005344998,0.0010647376,0.0010164691,0.00025385714],"domain_scores_gemma":[0.9918835,0.0043950006,0.0005848984,0.0011187532,0.0017412121,0.0002765221],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042351712,0.00090821105,0.0009089654,0.0067146393,0.0021643909,0.0036008167,0.002650202,0.0022810823,0.0041306913],"category_scores_gemma":[0.016800905,0.0008170288,0.001973252,0.0038837737,0.0010940232,0.008107404,0.00330547,0.0023165364,0.0033296915],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007791155,0.00047830498,0.005772696,0.00195384,0.00060455024,0.001110167,0.003654721,0.007381264,0.031213336,0.09033408,0.054286364,0.8024316],"study_design_scores_gemma":[0.00018081255,0.00032051333,0.0066335923,0.00049907906,0.0008971352,0.0024336095,0.0038072052,0.51111746,0.055483043,0.23601542,0.18228967,0.00032251806],"about_ca_topic_score_codex":0.0063974434,"about_ca_topic_score_gemma":0.008842284,"teacher_disagreement_score":0.0067146393,"about_ca_system_score_codex":0.0011089076,"about_ca_system_score_gemma":0.0024873174,"threshold_uncertainty_score":0.022398055},"labels":[],"label_agreement":null},{"id":"W2397043688","doi":"","title":"Strong Baselines for Cross-Lingual Entity Linking.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Entity linking; Simplicity; Task (project management); Context (archaeology); Natural language processing; Artificial intelligence; Competition (biology); Population; Geography; Knowledge base; Medicine","score_opus":0.029749906330027324,"score_gpt":0.2909864075180485,"score_spread":0.26123650118802116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397043688","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041409627,0.048612755,0.7387895,0.004361502,0.0052409475,0.0017265587,0.03229026,0.066935755,0.060633108],"genre_scores_gemma":[0.21966568,0.0075449585,0.57649434,0.0035813833,0.002064073,0.0025891224,0.14936778,0.0054289107,0.03326377],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98058367,0.0070795654,0.001426765,0.0045663947,0.0052146832,0.0011289273],"domain_scores_gemma":[0.96502995,0.012790507,0.0008868546,0.014169883,0.006124049,0.0009987443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02233381,0.004418817,0.0025138154,0.012019241,0.0052508777,0.0068273586,0.008131956,0.0060464703,0.023300523],"category_scores_gemma":[0.04530772,0.0015869372,0.0030628678,0.010146371,0.0015454192,0.015626388,0.012599599,0.0063953954,0.024167031],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014087361,0.0014235514,0.005762871,0.0026258035,0.0016924399,0.0004933699,0.00039508473,0.018709764,0.011076606,0.021666925,0.23670007,0.6980448],"study_design_scores_gemma":[0.0008647451,0.0013208671,0.015296603,0.0010799917,0.0019657207,0.0030026103,0.0015047574,0.43598202,0.045006227,0.1416947,0.3517319,0.0005498295],"about_ca_topic_score_codex":0.0053544743,"about_ca_topic_score_gemma":0.014183195,"teacher_disagreement_score":0.023300523,"about_ca_system_score_codex":0.0018836233,"about_ca_system_score_gemma":0.0029796711,"threshold_uncertainty_score":0.118113935},"labels":[],"label_agreement":null},{"id":"W2397558838","doi":"","title":"Domino: SAIC's English Entity-Linking System.","year":2012,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Domino; Supervisor; Competition (biology); Task (project management); Computer science; Baseline (sea); Domino effect; Plan (archaeology); Engineering; Political science","score_opus":0.009458581956566927,"score_gpt":0.22989732797939072,"score_spread":0.2204387460228238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397558838","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015400331,0.0025347301,0.13681075,0.0018121994,0.00082541525,0.0012284538,0.112071745,0.62529296,0.104023375],"genre_scores_gemma":[0.086558856,0.0013329502,0.27855444,0.0017381499,0.00030169135,0.0013200914,0.5069749,0.05652851,0.06669039],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977754,0.00040103533,0.0002744582,0.0006465482,0.000745119,0.00015753368],"domain_scores_gemma":[0.9946701,0.0012427026,0.00039321015,0.0019445018,0.0013337034,0.000415816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039250096,0.002004781,0.0013150921,0.006044691,0.0016971381,0.004843361,0.0034447485,0.0014514419,0.054804437],"category_scores_gemma":[0.014249467,0.0016659732,0.0009817166,0.0039261472,0.0007766292,0.009796556,0.006455677,0.002003034,0.04770382],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063754304,0.00013122773,0.002948322,0.0010729608,0.00012109528,0.0005096836,0.0006886029,0.00088322675,0.0060447627,0.010352896,0.8296108,0.14699884],"study_design_scores_gemma":[0.00026268806,0.00008276559,0.0030588422,0.00024933266,0.000065598724,0.000609397,0.0003532785,0.01678634,0.011251999,0.007018185,0.96013176,0.00012970116],"about_ca_topic_score_codex":0.012076184,"about_ca_topic_score_gemma":0.015890239,"teacher_disagreement_score":0.054804437,"about_ca_system_score_codex":0.0016812795,"about_ca_system_score_gemma":0.0030653789,"threshold_uncertainty_score":0.18333912},"labels":[],"label_agreement":null},{"id":"W2397624501","doi":"","title":"BUPTTeam Participation at TAC 2011 Recognizing Textual Entailment.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Textual entailment; Logical consequence; Task (project management); Computer science; Natural language processing; Similarity (geometry); Artificial intelligence; Information retrieval; Engineering","score_opus":0.033091363838856074,"score_gpt":0.2625577714443959,"score_spread":0.22946640760553982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397624501","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36075148,0.004633158,0.31909114,0.010749274,0.0075521795,0.0077824057,0.04173202,0.07689706,0.17081131],"genre_scores_gemma":[0.49063474,0.00076560804,0.24983415,0.0023302345,0.0011733375,0.0045949016,0.11946294,0.0047251172,0.12647893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99095726,0.00403577,0.00042387802,0.0012914285,0.002595649,0.0006960086],"domain_scores_gemma":[0.98660487,0.0056558475,0.00019061133,0.0024378423,0.0037023532,0.0014084298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010640823,0.0017928317,0.00190003,0.0023382595,0.0025571052,0.002975368,0.0037466502,0.0035792515,0.022265472],"category_scores_gemma":[0.027021084,0.0007279481,0.0010258367,0.0010242381,0.00081482145,0.0044861566,0.004749877,0.0030602028,0.014812709],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032746952,0.0030484186,0.004912207,0.0010670709,0.00037115067,0.0013999802,0.0033723712,0.0039141965,0.04717899,0.0037104045,0.3234378,0.6043127],"study_design_scores_gemma":[0.0016475179,0.00500316,0.026908815,0.0002855679,0.0004063471,0.0036425935,0.0033066187,0.16925012,0.12732942,0.009752731,0.65192246,0.00054456363],"about_ca_topic_score_codex":0.009223177,"about_ca_topic_score_gemma":0.012204561,"teacher_disagreement_score":0.022265472,"about_ca_system_score_codex":0.0013813581,"about_ca_system_score_gemma":0.0021044952,"threshold_uncertainty_score":0.07448542},"labels":[],"label_agreement":null},{"id":"W2397665446","doi":"10.3233/978-1-61499-289-9-594","title":"Engineering Natural Language Processing Solutions for Structured Information from Clinical Text: Extracting Sentinel Events from Palliative Care Consult Letters","year":2013,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Information extraction; Terminology; Artificial intelligence; Natural language processing; Natural language; SNOMED CT; Data extraction; Event (particle physics); Information retrieval; Bridge (graph theory); Machine learning; Medicine; MEDLINE; Linguistics","score_opus":0.043253344626875835,"score_gpt":0.3589841908514617,"score_spread":0.31573084622458586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397665446","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041916873,0.0006359608,0.9427845,0.0016265666,0.0001172625,0.0014731785,0.0069361,0.002636536,0.0018730981],"genre_scores_gemma":[0.069628306,0.0004745326,0.9178529,0.00015203017,0.00008440106,0.0008149285,0.010165577,0.00016193655,0.00066545996],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9937395,0.0027109683,0.0012527282,0.00094278593,0.00117796,0.00017602097],"domain_scores_gemma":[0.9658712,0.026502525,0.0024175523,0.0014489302,0.0035168685,0.00024294895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057068784,0.001183517,0.0008886238,0.007782974,0.0012656003,0.0033956168,0.0011159297,0.0012607067,0.0020471283],"category_scores_gemma":[0.025305672,0.00057439785,0.001673965,0.005244376,0.000936418,0.0032678933,0.0022773317,0.0013865919,0.0021516657],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043881164,0.00039486252,0.013832627,0.005682822,0.00024856592,0.0040390254,0.010533454,0.030592607,0.045199,0.020358326,0.018709298,0.8499707],"study_design_scores_gemma":[0.00027658418,0.0005361395,0.017260192,0.0012985104,0.0006415149,0.004539319,0.017318975,0.586567,0.08809356,0.12848672,0.15461375,0.0003678002],"about_ca_topic_score_codex":0.0028990088,"about_ca_topic_score_gemma":0.0033983092,"teacher_disagreement_score":0.007782974,"about_ca_system_score_codex":0.0011318739,"about_ca_system_score_gemma":0.003971802,"threshold_uncertainty_score":0.03018123},"labels":[],"label_agreement":null},{"id":"W2397770075","doi":"","title":"Sagan in TAC2009: Using Support Vector Machines in Recognizing Textual Entailment and TE Search Pilot task","year":2009,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Textual entailment; Logical consequence; Computer science; Artificial intelligence; Natural language processing; Pascal (unit); Classifier (UML); Support vector machine; Set (abstract data type); Task (project management); Semantic similarity; Programming language","score_opus":0.02585380174903698,"score_gpt":0.2879791380683192,"score_spread":0.26212533631928225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397770075","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34951776,0.0033553538,0.4703686,0.0016045279,0.0007436649,0.0014798404,0.012868023,0.14121018,0.018851995],"genre_scores_gemma":[0.4838428,0.0005004282,0.4589411,0.0007525769,0.00019662947,0.0005628204,0.04495207,0.0010260886,0.009225467],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835676,0.00058145914,0.00016071185,0.00033634467,0.0004395099,0.00012522435],"domain_scores_gemma":[0.99760604,0.0010332816,0.00015227072,0.00046248108,0.0006133246,0.00013260634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003175744,0.0011081753,0.0011671704,0.0022113465,0.0008411823,0.0014806483,0.0017970453,0.0019388333,0.0038467369],"category_scores_gemma":[0.007956866,0.00033018892,0.00064157444,0.0015041333,0.00039542723,0.0031217278,0.0012972159,0.0013873429,0.0034019025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001627639,0.000878558,0.006408363,0.0004603923,0.00022497165,0.00061232917,0.00040162564,0.013944086,0.025838152,0.003382874,0.0939145,0.85230654],"study_design_scores_gemma":[0.0003490005,0.001030581,0.008449531,0.000081168066,0.00025484207,0.0010590992,0.00051976973,0.8455348,0.056150176,0.013569062,0.07286111,0.00014099805],"about_ca_topic_score_codex":0.009291537,"about_ca_topic_score_gemma":0.012577917,"teacher_disagreement_score":0.009291537,"about_ca_system_score_codex":0.00067363336,"about_ca_system_score_gemma":0.0010166309,"threshold_uncertainty_score":0.018474877},"labels":[],"label_agreement":null},{"id":"W2398419065","doi":"","title":"Characterizing In-text Citations using N-gram Distributions","year":2015,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Citation; Rhetorical question; Computer science; Information retrieval; Metadata; Natural language processing; Linguistics; Citation analysis; Artificial intelligence; World Wide Web; Philosophy","score_opus":0.042747938807201326,"score_gpt":0.2711949104622968,"score_spread":0.2284469716550955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398419065","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41847724,0.009979811,0.53771394,0.0057012737,0.001245199,0.00027318537,0.009282138,0.0054395683,0.011887637],"genre_scores_gemma":[0.9425687,0.003028095,0.036278423,0.00040916767,0.0028162948,0.00020343196,0.00676366,0.00050025934,0.0074320696],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995865,0.0016885012,0.00037451554,0.0007517153,0.0009346995,0.000385595],"domain_scores_gemma":[0.9434935,0.045217182,0.0034363251,0.0023495352,0.004208138,0.0012953384],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0039817216,0.0012527895,0.001996689,0.011729612,0.0020061915,0.0046700872,0.0012495719,0.0030995598,0.0046492126],"category_scores_gemma":[0.04828792,0.00051515986,0.0011058316,0.011714205,0.0010237937,0.0068253158,0.0020873295,0.0020942988,0.004080205],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002230546,0.0011744376,0.10357068,0.0025844465,0.0008684678,0.0024225882,0.002483007,0.098171614,0.04715884,0.087746255,0.0604879,0.59110135],"study_design_scores_gemma":[0.00008412635,0.00027287577,0.021877844,0.0001884784,0.0002394426,0.0010744785,0.0005421806,0.81990653,0.010032371,0.13073462,0.01493268,0.00011435414],"about_ca_topic_score_codex":0.0024283426,"about_ca_topic_score_gemma":0.0038314862,"teacher_disagreement_score":0.9960183,"about_ca_system_score_codex":0.0013558116,"about_ca_system_score_gemma":0.001594706,"threshold_uncertainty_score":0.021057606},"labels":[],"label_agreement":null},{"id":"W2398735901","doi":"","title":"The Influence of Contextual Variability in Word Learning.","year":2014,"lang":"en","type":"article","venue":"Grantee Submission","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Optimal distinctiveness theory; Word (group theory); Computer science; Natural language processing; Word lists by frequency; Context (archaeology); Artificial intelligence; Newspaper; Semantic similarity; Linguistics; Psychology; History; Social psychology","score_opus":0.010629203589932357,"score_gpt":0.23268368924890231,"score_spread":0.22205448565896996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398735901","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.876035,0.003341045,0.11095843,0.00054380053,0.00014950678,0.0001462711,0.0014880559,0.0008339344,0.006503908],"genre_scores_gemma":[0.98949414,0.00019672394,0.008566943,0.00007618666,0.000045593944,0.000053842712,0.0010863993,0.00015736339,0.00032269847],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99109113,0.0041213054,0.00052614685,0.0034025277,0.000619332,0.00023961642],"domain_scores_gemma":[0.9143174,0.06806248,0.003920926,0.010664943,0.0018867643,0.0011474704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010390741,0.0006538059,0.000816761,0.0014732716,0.00092852995,0.0030021085,0.000961081,0.00085012737,0.0018613329],"category_scores_gemma":[0.08065256,0.0005086261,0.0012301421,0.0015953348,0.0018798107,0.0040945997,0.0030275541,0.0023176253,0.0005784693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020375843,0.00038787583,0.6254486,0.00090416346,0.0018178681,0.00061418675,0.0042744237,0.032276805,0.047947746,0.011365187,0.004545485,0.2683801],"study_design_scores_gemma":[0.00012373296,0.0009286799,0.67479426,0.00020612852,0.0010120315,0.0021764692,0.0019243318,0.22036931,0.022330046,0.06431657,0.011575401,0.00024307778],"about_ca_topic_score_codex":0.0030217636,"about_ca_topic_score_gemma":0.004597002,"teacher_disagreement_score":0.010390741,"about_ca_system_score_codex":0.00065912475,"about_ca_system_score_gemma":0.00060782186,"threshold_uncertainty_score":0.054952145},"labels":[],"label_agreement":null},{"id":"W2398787226","doi":"","title":"A Neurally Plausible Encoding of Word Order Information into a Semantic Vector Space","year":2013,"lang":"en","type":"article","venue":"eScholarship (California Digital Library)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Encoding (memory); Word (group theory); Computer science; Natural language processing; Semantics (computer science); Distributional semantics; Vector space; Artificial intelligence; Word order; Space (punctuation); Vector space model; Theoretical computer science; Semantic similarity; Mathematics","score_opus":0.011266136361472586,"score_gpt":0.20106291797266929,"score_spread":0.1897967816111967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398787226","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020640695,0.00014916433,0.9744486,0.00060118677,0.000070333626,0.000020594976,0.00018123674,0.0004633687,0.0034248254],"genre_scores_gemma":[0.5749284,0.00048518667,0.41661537,0.00019974768,0.00013510333,0.0000797591,0.00045577745,0.00011194308,0.006988684],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99970526,0.00009938095,0.000018816245,0.00009339174,0.00006137374,0.000021699667],"domain_scores_gemma":[0.99947435,0.00024850623,0.000049554576,0.000133829,0.00006805189,0.000025699577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006273674,0.00033584397,0.00038091553,0.00055001525,0.00032349076,0.0015440146,0.0011207392,0.00077597017,0.0054129395],"category_scores_gemma":[0.003659262,0.0003276499,0.00060722633,0.0007834619,0.00090173323,0.004668485,0.0008038936,0.0013409038,0.0008603697],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016917642,0.00011214964,0.0013653141,0.00020139056,0.00008858866,0.00014310818,0.00029666,0.11459024,0.020976184,0.55619013,0.0047784895,0.30108866],"study_design_scores_gemma":[0.000018443285,0.00006130952,0.00057523686,0.000028832643,0.000019960897,0.00012832686,0.000052893356,0.67035615,0.0049167858,0.32021564,0.0036003515,0.000026103678],"about_ca_topic_score_codex":0.0014814389,"about_ca_topic_score_gemma":0.0025541843,"teacher_disagreement_score":0.0054129395,"about_ca_system_score_codex":0.00064179336,"about_ca_system_score_gemma":0.0005556329,"threshold_uncertainty_score":0.01810807},"labels":[],"label_agreement":null},{"id":"W2399031877","doi":"","title":"Université de Montréal at the NTCIR-11 IMine Task","year":2014,"lang":"fr","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ranking (information retrieval); Computer science; Task (project management); Set (abstract data type); Information retrieval; Embedding; Resource (disambiguation); Artificial intelligence","score_opus":0.012377848107241953,"score_gpt":0.17961887650148303,"score_spread":0.16724102839424107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399031877","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10844115,0.017056886,0.03370252,0.01857771,0.009770625,0.005371567,0.44932133,0.036292415,0.3214659],"genre_scores_gemma":[0.20806938,0.0021679346,0.052099,0.0035138316,0.001248013,0.0027517113,0.55379117,0.003273685,0.17308526],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9955426,0.0013905448,0.00013922241,0.00090219517,0.0011656272,0.00085982017],"domain_scores_gemma":[0.991398,0.0016195202,0.00021009038,0.0015478756,0.0030887914,0.0021357688],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061214785,0.0033575392,0.00277071,0.002404542,0.0040629325,0.0042160302,0.0021803665,0.0020658928,0.050473686],"category_scores_gemma":[0.011956557,0.00067262555,0.00113302,0.0023629013,0.00085131935,0.003333081,0.0034981654,0.0026412383,0.032863077],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008679324,0.00042212542,0.0023210712,0.00058579794,0.00012680482,0.0003891196,0.0003859089,0.0013172688,0.0041931085,0.0021318113,0.9258318,0.061427295],"study_design_scores_gemma":[0.0011058418,0.0007625046,0.025086172,0.00031409762,0.00019824818,0.0007002095,0.0013337918,0.019179383,0.011067873,0.0057913046,0.9341677,0.00029294487],"about_ca_topic_score_codex":0.33288407,"about_ca_topic_score_gemma":0.37761837,"teacher_disagreement_score":0.33288407,"about_ca_system_score_codex":0.0055854954,"about_ca_system_score_gemma":0.009039883,"threshold_uncertainty_score":0.66189295},"labels":[],"label_agreement":null},{"id":"W2399435498","doi":"","title":"Recognizing Textual Entailment with Logical Inference.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Logical consequence; Textual entailment; Natural language processing; Computer science; Paraphrase; WordNet; Artificial intelligence; Sentence; Inference; Task (project management); Interpretation (philosophy); Programming language","score_opus":0.027330417824150494,"score_gpt":0.25628214586897535,"score_spread":0.22895172804482486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399435498","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020757427,0.0006583228,0.9491832,0.0012090387,0.00011030338,0.0005602394,0.001697926,0.019497022,0.0063263727],"genre_scores_gemma":[0.15406619,0.00048761256,0.83571106,0.00040974838,0.00013974264,0.00022889777,0.005737798,0.00046889618,0.0027500195],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99522865,0.0018443001,0.0004857441,0.0011299092,0.001112382,0.00019910464],"domain_scores_gemma":[0.976553,0.01660712,0.001829484,0.0024461877,0.0022202597,0.0003439955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007153747,0.0014718932,0.00091457274,0.003916665,0.0014760386,0.004461303,0.0029845554,0.0019560552,0.011838146],"category_scores_gemma":[0.03890451,0.00092757185,0.0025311594,0.002003979,0.0013983853,0.010437475,0.0037632838,0.002266412,0.005726727],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009507514,0.00050237414,0.013561949,0.0026772,0.00057518657,0.0015113476,0.0037225499,0.015408521,0.034624014,0.060838155,0.04656724,0.81906074],"study_design_scores_gemma":[0.00017983916,0.00033853273,0.005724277,0.00040196042,0.0006234026,0.002515204,0.0024790585,0.6163823,0.08442607,0.21974349,0.067028716,0.00015716905],"about_ca_topic_score_codex":0.0031155874,"about_ca_topic_score_gemma":0.005082103,"teacher_disagreement_score":0.011838146,"about_ca_system_score_codex":0.001289451,"about_ca_system_score_gemma":0.002047667,"threshold_uncertainty_score":0.039602518},"labels":[],"label_agreement":null},{"id":"W2399456070","doi":"10.21437/interspeech.2013-596","title":"Investigation of recurrent-neural-network architectures and learning methods for spoken language understanding","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":432,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Spoken language; Artificial neural network; Natural language processing; Recurrent neural network","score_opus":0.0788804892206361,"score_gpt":0.32925922817297804,"score_spread":0.2503787389523419,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399456070","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09572335,0.002678476,0.89152884,0.00080031354,0.00011591896,0.000095425115,0.00029303832,0.0037008403,0.005063816],"genre_scores_gemma":[0.68622357,0.0015070059,0.30507758,0.00021227569,0.00008360602,0.00013528162,0.0009444337,0.00028103674,0.005535188],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993325,0.00025816893,0.00003857666,0.00019068344,0.00011906815,0.000060991093],"domain_scores_gemma":[0.99699616,0.0018439846,0.00016056233,0.00028898657,0.00063363457,0.0000767626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038715878,0.0009838898,0.00062908704,0.00066422444,0.00033912846,0.001034648,0.0016410556,0.0009728762,0.00223805],"category_scores_gemma":[0.0075877476,0.00048507098,0.00057955924,0.0006231792,0.00039059477,0.0035164107,0.00068934314,0.0018481163,0.0007099769],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022029666,0.00021481191,0.0022182649,0.0002376101,0.00018000815,0.000097116004,0.00026770003,0.6202718,0.010889617,0.021198735,0.002546108,0.34165788],"study_design_scores_gemma":[0.000004346771,0.000035754765,0.00013380438,0.000009121521,0.000012025272,0.000010380292,0.000014114221,0.9947379,0.0017258788,0.0028856029,0.00042575374,0.0000053123986],"about_ca_topic_score_codex":0.007121918,"about_ca_topic_score_gemma":0.01051028,"teacher_disagreement_score":0.007121918,"about_ca_system_score_codex":0.00094961876,"about_ca_system_score_gemma":0.00095068914,"threshold_uncertainty_score":0.02047515},"labels":[],"label_agreement":null},{"id":"W2399623965","doi":"","title":"Generate Compressed Sentences with Stanford Typed Dependencies towards Abstractive Summarization.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Natural language processing; Ranking (information retrieval); Natural language generation; Artificial intelligence; Selection (genetic algorithm); Metric (unit); Process (computing); Natural language; Programming language","score_opus":0.023241655905932915,"score_gpt":0.2320686903214139,"score_spread":0.20882703441548098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399623965","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02128097,0.00039321222,0.9637286,0.00046381945,0.00022732021,0.00044822536,0.0026644284,0.007988401,0.0028050388],"genre_scores_gemma":[0.11319832,0.00028230788,0.8697649,0.00022321955,0.00016664092,0.00048274457,0.010479797,0.0006149267,0.0047872225],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99893206,0.0004652219,0.00008886548,0.00019839813,0.0002754911,0.000040024606],"domain_scores_gemma":[0.9964684,0.0016243001,0.0003069666,0.00055070495,0.00096092105,0.00008874444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012590627,0.001160545,0.0005227059,0.0012741879,0.0004202514,0.0008275518,0.0008984062,0.00066454295,0.005762292],"category_scores_gemma":[0.0064316266,0.00029191002,0.0006592357,0.0009533987,0.00027535707,0.0014544396,0.0008989199,0.000903362,0.0031816107],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004790437,0.00039422014,0.0017259801,0.0013051224,0.00021485605,0.0006801297,0.0010235937,0.031548306,0.14643288,0.019594895,0.054427255,0.7421738],"study_design_scores_gemma":[0.00017519505,0.000847843,0.0033081402,0.00013221288,0.00032403175,0.00074564025,0.0005618812,0.6329046,0.2144612,0.05449926,0.09189116,0.0001488047],"about_ca_topic_score_codex":0.00084785855,"about_ca_topic_score_gemma":0.00192276,"teacher_disagreement_score":0.005762292,"about_ca_system_score_codex":0.00032560466,"about_ca_system_score_gemma":0.0008044005,"threshold_uncertainty_score":0.019276738},"labels":[],"label_agreement":null},{"id":"W2399880602","doi":"10.1609/aaai.v31i1.10983","title":"A Hierarchical Latent Variable Encoder-Decoder Model for Generating Dialogues","year":2017,"lang":"en","type":"preprint","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":266,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; McGill University; Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Samsung; Compute Canada; Samsung Advanced Institute of Technology; Canadian Institute for Advanced Research","keywords":"Computer science; Latent variable; Generative grammar; Generative model; Latent variable model; Artificial intelligence; Context (archaeology); Variable (mathematics); Artificial neural network; Process (computing); Encoder; Machine learning; Task (project management); Mathematics; Engineering","score_opus":0.21650778029913978,"score_gpt":0.3321390255953345,"score_spread":0.1156312452961947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399880602","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015738621,0.00027063556,0.98024726,0.0004614954,0.000065031985,0.000073460884,0.00034970941,0.0011700786,0.0016237025],"genre_scores_gemma":[0.62573504,0.00037100888,0.36162838,0.00032088428,0.00014098587,0.00054607977,0.0012538747,0.00034758492,0.009656114],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991054,0.00045487037,0.000039564453,0.00021828999,0.0001127911,0.00006908795],"domain_scores_gemma":[0.99777204,0.0016767952,0.00011540562,0.00014617297,0.00021177533,0.000077816905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018265644,0.000785258,0.0006308528,0.00066649966,0.0004003156,0.0008479111,0.00156732,0.0013519607,0.004114923],"category_scores_gemma":[0.0058167074,0.0006030776,0.00086720847,0.000783503,0.0006976786,0.0014624656,0.0010498702,0.0020712218,0.001402666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000594786,0.00019414887,0.0021408552,0.00026408187,0.00013910637,0.00031589676,0.0008445431,0.7428105,0.010041214,0.08645752,0.005839127,0.15035823],"study_design_scores_gemma":[0.000014451246,0.00001772841,0.00007847111,0.000005349767,0.000009764588,0.000023355993,0.000008004518,0.9899182,0.0004981742,0.008949878,0.00047022421,0.0000064119804],"about_ca_topic_score_codex":0.0048897774,"about_ca_topic_score_gemma":0.007904901,"teacher_disagreement_score":0.0048897774,"about_ca_system_score_codex":0.0009869706,"about_ca_system_score_gemma":0.0011937665,"threshold_uncertainty_score":0.013765812},"labels":[],"label_agreement":null},{"id":"W2399977860","doi":"","title":"TAC 2008 Question Answering Experiments at Tokyo Institute of Technology.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Question answering; Information retrieval","score_opus":0.01433929541817701,"score_gpt":0.25980218744265965,"score_spread":0.24546289202448263,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399977860","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5665834,0.0030089377,0.070813954,0.015984738,0.0057706404,0.0066230698,0.12217591,0.03891582,0.17012352],"genre_scores_gemma":[0.62593406,0.00066150806,0.08931221,0.0030269374,0.0010831679,0.0068599875,0.18213421,0.0021269356,0.08886108],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9970867,0.001320959,0.00018232434,0.0007307388,0.00046142904,0.00021786556],"domain_scores_gemma":[0.9907584,0.004182876,0.00020681082,0.0014826315,0.0024772005,0.00089206436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045646466,0.0009946461,0.001286379,0.0010591975,0.0018566118,0.001871238,0.0017677743,0.0023076544,0.029423196],"category_scores_gemma":[0.011363765,0.00052770047,0.0005618366,0.0011400975,0.00053179497,0.0037574978,0.0013535176,0.0025401928,0.010552635],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003832119,0.004927693,0.0049193394,0.0006713272,0.00018255491,0.00048104013,0.0017464664,0.0016859844,0.023545599,0.005908508,0.82896036,0.12313909],"study_design_scores_gemma":[0.006643027,0.00469081,0.053554185,0.00018427984,0.00057669147,0.0012335523,0.0019940278,0.061030168,0.041136615,0.019699674,0.80889636,0.00036069163],"about_ca_topic_score_codex":0.008386229,"about_ca_topic_score_gemma":0.0102646,"teacher_disagreement_score":0.029423196,"about_ca_system_score_codex":0.001075007,"about_ca_system_score_gemma":0.0016426507,"threshold_uncertainty_score":0.098430455},"labels":[],"label_agreement":null},{"id":"W2400124070","doi":"","title":"Decayed DivRank for Guided Summarization.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Relevance (law); Sentence; Novelty; Multi-document summarization; Ranking (information retrieval); Information retrieval; Artificial intelligence; Natural language processing","score_opus":0.029807933303666274,"score_gpt":0.2607443788628585,"score_spread":0.23093644555919224,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2400124070","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004200247,0.0008077883,0.99186796,0.00014140653,0.00008088058,0.000110249515,0.00020330989,0.0012585807,0.0013297452],"genre_scores_gemma":[0.17319144,0.0009540574,0.81570244,0.00026281443,0.0003947845,0.0004183866,0.0019089229,0.00045685525,0.006710319],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977756,0.0009386664,0.00016705865,0.00043351224,0.0005945448,0.00009043279],"domain_scores_gemma":[0.9957495,0.0019124394,0.0004614303,0.0007764409,0.00095995853,0.00014025238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019162166,0.0011473174,0.001051716,0.0017835118,0.00068010343,0.0014783272,0.0013337613,0.00097522035,0.0030738672],"category_scores_gemma":[0.008339145,0.0003620003,0.0007035945,0.0017611913,0.0006728137,0.0026429547,0.0011068786,0.0012844016,0.002009006],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034157984,0.00015292809,0.00130473,0.0007261563,0.00016945152,0.00023391182,0.0003930639,0.08767623,0.030301468,0.07763488,0.021055616,0.78001],"study_design_scores_gemma":[0.00006693318,0.00029610732,0.0005137099,0.000049088918,0.00007552584,0.00023254301,0.0000789763,0.8703524,0.016936833,0.0784809,0.032859,0.00005796026],"about_ca_topic_score_codex":0.0012216538,"about_ca_topic_score_gemma":0.0030873076,"teacher_disagreement_score":0.0030738672,"about_ca_system_score_codex":0.0006950486,"about_ca_system_score_gemma":0.0010134199,"threshold_uncertainty_score":0.0102831125},"labels":[],"label_agreement":null},{"id":"W2400166611","doi":"","title":"Supervised Learning for Linking Named Entities to Knowledge Base Entries.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Knowledge base; Entity linking; Task (project management); Artificial intelligence; Rank (graph theory); Base (topology); Information retrieval; Range (aeronautics); Machine learning; Supervised learning; Feature (linguistics); Simple (philosophy); Data mining; Natural language processing; Mathematics; Artificial neural network","score_opus":0.02576742071232573,"score_gpt":0.252675748689076,"score_spread":0.22690832797675026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2400166611","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008009322,0.0010464651,0.986334,0.00022867757,0.00008214778,0.0001960136,0.00063803926,0.0022678236,0.0011975999],"genre_scores_gemma":[0.16174877,0.0011116972,0.8251028,0.00022806856,0.00032267338,0.00057904294,0.008093078,0.00023011395,0.0025837887],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.995282,0.0022317462,0.00029602583,0.0010652046,0.0009960239,0.00012901146],"domain_scores_gemma":[0.98344374,0.010809165,0.0015339571,0.002572848,0.001397212,0.00024314292],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004634102,0.0012834491,0.001307284,0.0057761185,0.0011446119,0.0019015652,0.0035681163,0.0019526287,0.002141329],"category_scores_gemma":[0.020996116,0.00066209957,0.001233966,0.0050513037,0.00090133765,0.005252574,0.002057442,0.0019422644,0.003116673],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000203742,0.0008089344,0.0054428,0.0010670382,0.00046439306,0.00019062952,0.00038586525,0.09476987,0.0053863814,0.018164461,0.020854682,0.8522612],"study_design_scores_gemma":[0.00003823404,0.00011867154,0.0013050764,0.00009196131,0.00009446281,0.0002291401,0.0001599787,0.9142372,0.007368287,0.06633348,0.009979313,0.000044213997],"about_ca_topic_score_codex":0.0017770318,"about_ca_topic_score_gemma":0.0033559084,"teacher_disagreement_score":0.0057761185,"about_ca_system_score_codex":0.0009102608,"about_ca_system_score_gemma":0.0015953137,"threshold_uncertainty_score":0.024507761},"labels":[],"label_agreement":null},{"id":"W2400208831","doi":"","title":"Bag of Senses Versus Bag of Words: Comparing Semantic and Lexical Approaches on Sentence Extraction.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Natural language processing; Computer science; Artificial intelligence; Sentence; Bag-of-words model","score_opus":0.06651501571551183,"score_gpt":0.27276171798514703,"score_spread":0.2062467022696352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2400208831","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21559256,0.024059506,0.71986663,0.0032614176,0.0012718389,0.0012215103,0.01192523,0.0048417943,0.01795942],"genre_scores_gemma":[0.5066038,0.00648223,0.46547997,0.0005303111,0.0008658673,0.0008284018,0.015986774,0.0005996234,0.0026232023],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99298877,0.004194082,0.00059551536,0.0006580477,0.0013631963,0.0002004563],"domain_scores_gemma":[0.9701173,0.02514009,0.00095201813,0.0009637014,0.0025323022,0.00029455492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007183796,0.0007949167,0.0011846862,0.008631789,0.0007236032,0.0028914667,0.0012086682,0.0010218862,0.0022094913],"category_scores_gemma":[0.036689963,0.000321314,0.0010954761,0.008057047,0.0006211335,0.007742796,0.001955882,0.0009964966,0.001570564],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016726247,0.00035580172,0.017998418,0.0024401385,0.0006234554,0.00014157513,0.0012050653,0.0033691048,0.009508065,0.01328695,0.022123449,0.9272753],"study_design_scores_gemma":[0.0008251234,0.0025071516,0.097096495,0.00223575,0.00360824,0.0026449035,0.010811001,0.49385586,0.04143739,0.24182038,0.10257254,0.0005852109],"about_ca_topic_score_codex":0.0017710018,"about_ca_topic_score_gemma":0.0041172076,"teacher_disagreement_score":0.008631789,"about_ca_system_score_codex":0.00061428576,"about_ca_system_score_gemma":0.0015607895,"threshold_uncertainty_score":0.03799194},"labels":[],"label_agreement":null},{"id":"W2400500583","doi":"","title":"Predicting Summary Quality using Limited Human Input.","year":2009,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Group cohesiveness; Pyramid (geometry); Similarity (geometry); Quality (philosophy); Data mining; Correlation; Artificial intelligence; Machine learning; Mathematics; Image (mathematics); Psychology","score_opus":0.03254081077730608,"score_gpt":0.30805292853572885,"score_spread":0.27551211775842277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2400500583","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8615314,0.0025785028,0.116916165,0.00038170506,0.00019694895,0.0006102838,0.005041939,0.008330552,0.00441246],"genre_scores_gemma":[0.9322561,0.0001955771,0.060450252,0.00006622345,0.000051832558,0.0001554395,0.005809923,0.00012537195,0.00088927004],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9921616,0.0040880903,0.0006561552,0.0016336626,0.0012217972,0.00023873868],"domain_scores_gemma":[0.9136767,0.067542516,0.004549368,0.006432476,0.0065052067,0.001293845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010406561,0.0016482939,0.0010518248,0.0026906852,0.0005435313,0.0018515262,0.0010043298,0.0016137541,0.0014135048],"category_scores_gemma":[0.06372171,0.00032856202,0.00074166345,0.0019140118,0.00036771293,0.0019832293,0.0009573755,0.0010622811,0.00095235003],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0048328307,0.0016444362,0.17972517,0.002605731,0.0018429135,0.0007021635,0.002316414,0.13329679,0.03691737,0.001136367,0.0153896995,0.61959016],"study_design_scores_gemma":[0.00034157408,0.0042141983,0.13017224,0.00015109754,0.0006170741,0.0004304011,0.0007586161,0.80893826,0.044659175,0.0030214868,0.0064450116,0.00025080092],"about_ca_topic_score_codex":0.0048107044,"about_ca_topic_score_gemma":0.0076718656,"teacher_disagreement_score":0.010406561,"about_ca_system_score_codex":0.0008886699,"about_ca_system_score_gemma":0.0006634033,"threshold_uncertainty_score":0.05503583},"labels":[],"label_agreement":null},{"id":"W2401064360","doi":"10.5281/zenodo.1417003","title":"Are Poetry And Lyrics All That Different?","year":2014,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Lyrics; Poetry; Newspaper; Set (abstract data type); Literature; Linguistics; Computer science; Natural language processing; Art; Artificial intelligence; Philosophy; Sociology","score_opus":0.05816569313429183,"score_gpt":0.23952274540260524,"score_spread":0.1813570522683134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401064360","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011006431,0.014205789,0.020796586,0.117603816,0.02725182,0.00005931677,0.0017287504,0.0006863011,0.80666125],"genre_scores_gemma":[0.5165054,0.01513937,0.015582399,0.024252774,0.028063465,0.00018251962,0.0039471583,0.002874858,0.3934521],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915576,0.00041119708,0.000044521486,0.00011218348,0.00020197302,0.00007432504],"domain_scores_gemma":[0.9981306,0.0005779882,0.0001452601,0.0002626004,0.0006823621,0.00020117096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00092655513,0.0004375809,0.0005617001,0.0015533749,0.0033144036,0.005983948,0.00069483405,0.0010608808,0.024792949],"category_scores_gemma":[0.0069572204,0.00026903916,0.00027807205,0.0021753544,0.0037450131,0.0077263056,0.0017586699,0.0030821064,0.011671742],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000088075,0.000019990393,0.00046382044,0.0001405617,0.00001674049,0.00007374609,0.00986038,0.000053423857,0.00036105202,0.5177861,0.4303615,0.040774506],"study_design_scores_gemma":[0.000007397943,0.000010189765,0.0011614338,0.000099836616,0.000008106428,0.00022308758,0.0042693666,0.00016340827,0.00026871727,0.0723796,0.9213887,0.000020105781],"about_ca_topic_score_codex":0.0038823904,"about_ca_topic_score_gemma":0.004786637,"teacher_disagreement_score":0.024792949,"about_ca_system_score_codex":0.0013006979,"about_ca_system_score_gemma":0.00068663945,"threshold_uncertainty_score":0.08294064},"labels":[],"label_agreement":null},{"id":"W2401350461","doi":"","title":"Learning Task Experiments in the TREC 2010 Legal Track.","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Computer science; Track (disk drive); Task (project management); Artificial intelligence; Natural language processing; Information retrieval; Machine learning; Engineering; Operating system","score_opus":0.021906744281011505,"score_gpt":0.2715438169130098,"score_spread":0.24963707263199827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401350461","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89176416,0.0038056287,0.025463479,0.0021104903,0.0015619074,0.012745491,0.025652578,0.0064115683,0.030484809],"genre_scores_gemma":[0.76946384,0.0013046435,0.09558403,0.003598262,0.001217559,0.01557549,0.08288879,0.0007888345,0.029578447],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9815134,0.009815409,0.0016101559,0.0021252516,0.003754562,0.0011812721],"domain_scores_gemma":[0.92601013,0.056474786,0.002337592,0.0059853937,0.006753121,0.0024389704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019719856,0.0034297134,0.0024818387,0.0020423953,0.0024548874,0.0020607852,0.0034047395,0.005121068,0.0059693814],"category_scores_gemma":[0.0584663,0.0010317626,0.0021526918,0.0024154133,0.0016847133,0.0050138277,0.0031256434,0.0052894666,0.0039975527],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.033861727,0.0942223,0.032652542,0.008133636,0.0030236416,0.001723427,0.004226071,0.10187259,0.03847426,0.0045123766,0.2630509,0.41424647],"study_design_scores_gemma":[0.04359426,0.072379306,0.1033854,0.00095942593,0.0029009588,0.0024665417,0.0044647306,0.5113902,0.10530081,0.018679675,0.13281433,0.0016642743],"about_ca_topic_score_codex":0.023160169,"about_ca_topic_score_gemma":0.02560101,"teacher_disagreement_score":0.023160169,"about_ca_system_score_codex":0.0028342018,"about_ca_system_score_gemma":0.003727158,"threshold_uncertainty_score":0.10428983},"labels":[],"label_agreement":null},{"id":"W2401480511","doi":"","title":"NUS at TAC 2008: Augumenting Timestamped Graphs with Event Information and Selectively Expanding Opinion Contexts.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Computer science; Physics; Astrophysics","score_opus":0.007914126589950565,"score_gpt":0.23338932594665995,"score_spread":0.2254751993567094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401480511","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.091778874,0.0043124184,0.8096427,0.004829449,0.0020231078,0.00045564133,0.03461125,0.035907608,0.016439022],"genre_scores_gemma":[0.52901804,0.000992202,0.39536557,0.0007876337,0.0010376151,0.0003596406,0.053228945,0.0024140521,0.016796378],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998461,0.0007049273,0.00007541854,0.0003147622,0.0003292042,0.00011459141],"domain_scores_gemma":[0.99468863,0.0024094086,0.00025006416,0.0015256617,0.0008473975,0.00027890847],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026545816,0.0009302221,0.00094708375,0.002698348,0.0009521546,0.00196023,0.0017010836,0.0013649423,0.0050942446],"category_scores_gemma":[0.01742178,0.00046318228,0.0010111309,0.0022437486,0.0005775223,0.0053101997,0.0026117582,0.0020398418,0.0032188531],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014289089,0.00040744376,0.009217609,0.00058856054,0.00029845975,0.0004588236,0.0016340284,0.035448633,0.009656866,0.06258467,0.36071065,0.51756525],"study_design_scores_gemma":[0.00020749801,0.00014669173,0.004039137,0.000105704865,0.00013851616,0.00026641798,0.0006513241,0.7263168,0.0070964997,0.17527524,0.085679,0.000077180084],"about_ca_topic_score_codex":0.010770943,"about_ca_topic_score_gemma":0.027688365,"teacher_disagreement_score":0.010770943,"about_ca_system_score_codex":0.0009965138,"about_ca_system_score_gemma":0.0011884788,"threshold_uncertainty_score":0.021416485},"labels":[],"label_agreement":null},{"id":"W2402510845","doi":"","title":"Overview of the TREC 2011 Legal Track.","year":2011,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Task (project management); Computer science; Track (disk drive); Rank (graph theory); Information retrieval; Natural language processing; Artificial intelligence","score_opus":0.1480039556696987,"score_gpt":0.28564236471141913,"score_spread":0.13763840904172042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402510845","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007077255,0.011218451,0.024273997,0.009921955,0.0043352116,0.0055919713,0.73966277,0.023936193,0.1739822],"genre_scores_gemma":[0.018830797,0.0046528494,0.03859375,0.0021082857,0.0012469525,0.0037468162,0.82291776,0.002837213,0.10506553],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934657,0.0016042899,0.0004904693,0.0007366166,0.0030650282,0.0006379708],"domain_scores_gemma":[0.98184574,0.0024783108,0.0008820911,0.0018958129,0.011193214,0.001704841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010583707,0.0017294128,0.0010844821,0.008979253,0.0020987,0.004006035,0.002149395,0.0016470733,0.107846335],"category_scores_gemma":[0.015358534,0.00085814187,0.0008404023,0.0095780855,0.00050122204,0.004801618,0.0016415234,0.0017674883,0.099614084],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007172587,0.00007410086,0.0003250593,0.00039578337,0.000013022003,0.000015889205,0.000040727675,0.0004922892,0.001152587,0.00073114986,0.95944524,0.03724253],"study_design_scores_gemma":[0.000111645146,0.00016315817,0.007964171,0.0002886725,0.000053314983,0.000109318484,0.00016347013,0.0024590387,0.0031381175,0.0019271815,0.9835012,0.000120743665],"about_ca_topic_score_codex":0.088469565,"about_ca_topic_score_gemma":0.115350276,"teacher_disagreement_score":0.107846335,"about_ca_system_score_codex":0.004322385,"about_ca_system_score_gemma":0.009431419,"threshold_uncertainty_score":0.3607819},"labels":[],"label_agreement":null},{"id":"W2402585298","doi":"","title":"Off to a cold start: New York University's 2013 knowledge base population systems","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Cold start (automotive); Base (topology); Population; Computer science; Knowledge base; Operations research; Mathematics; Engineering; Demography; Artificial intelligence; Sociology","score_opus":0.01512373262703934,"score_gpt":0.21910084551242845,"score_spread":0.2039771128853891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402585298","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4476658,0.004291824,0.2912805,0.04888041,0.006257331,0.0037733852,0.055138156,0.096298225,0.046414386],"genre_scores_gemma":[0.389488,0.0008311275,0.43050253,0.003553092,0.0008759419,0.0030842992,0.118382275,0.010923299,0.04235942],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9861898,0.0049091647,0.00091481203,0.0025396124,0.0047199405,0.0007268084],"domain_scores_gemma":[0.9630652,0.015778035,0.00041190055,0.0069826106,0.011590318,0.0021719332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030634241,0.00061905343,0.0017414412,0.0029093844,0.0035274439,0.004987887,0.0031426344,0.0017034294,0.007915934],"category_scores_gemma":[0.061477255,0.0011789281,0.001019141,0.0046289554,0.0014199602,0.0065185484,0.0043739043,0.004182661,0.0027379808],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018679929,0.00078240567,0.015764618,0.0004490913,0.00046814128,0.00029281317,0.004678006,0.02376922,0.004759388,0.022441855,0.5093929,0.4153336],"study_design_scores_gemma":[0.0020486992,0.00072252506,0.044213768,0.00028002277,0.00039430923,0.00023782205,0.004602942,0.24595226,0.025755916,0.031037433,0.644152,0.0006023701],"about_ca_topic_score_codex":0.18817154,"about_ca_topic_score_gemma":0.2548998,"teacher_disagreement_score":0.18817154,"about_ca_system_score_codex":0.007971333,"about_ca_system_score_gemma":0.009949721,"threshold_uncertainty_score":0.37415248},"labels":[],"label_agreement":null},{"id":"W2402842738","doi":"","title":"Stanford's Distantly-Supervised Slot-Filling System","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Coreference; Knowledge base; Inference; Task (project management); Process (computing); Entity linking; Information retrieval; Population; Natural language processing; Artificial intelligence; World Wide Web; Resolution (logic); Programming language","score_opus":0.019186594191083275,"score_gpt":0.22350375731268307,"score_spread":0.2043171631215998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402842738","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043174494,0.0006072609,0.6838678,0.00055080774,0.00029384144,0.00050044304,0.022169521,0.228436,0.02039977],"genre_scores_gemma":[0.1971111,0.0002221169,0.7334412,0.00039472256,0.00013642294,0.0005481233,0.04506729,0.0037700161,0.019309033],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990695,0.00016275619,0.00008699738,0.00036955575,0.00025170407,0.00005958409],"domain_scores_gemma":[0.99822587,0.00075804285,0.00008238257,0.0004096864,0.0004362387,0.000087805915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015629021,0.0007543392,0.0007425475,0.0015077306,0.0008799415,0.0008580355,0.0023708586,0.0008671647,0.018753063],"category_scores_gemma":[0.0036470853,0.0005860589,0.0007813676,0.0013248363,0.00036069422,0.0027108353,0.001715267,0.0009975624,0.009327877],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067224115,0.00027576304,0.003596239,0.000709719,0.00008827122,0.0003863737,0.00069336995,0.009278564,0.031114632,0.014337202,0.2743106,0.6645371],"study_design_scores_gemma":[0.00031386057,0.00030647218,0.008349207,0.000105944564,0.0001769054,0.0009979972,0.0003952584,0.5438654,0.08633876,0.032639336,0.32626164,0.00024922407],"about_ca_topic_score_codex":0.0050265943,"about_ca_topic_score_gemma":0.007027759,"teacher_disagreement_score":0.018753063,"about_ca_system_score_codex":0.0005897856,"about_ca_system_score_gemma":0.0015026267,"threshold_uncertainty_score":0.0627352},"labels":[],"label_agreement":null},{"id":"W2403788938","doi":"","title":"Using Unsupervised System with least linguistic features for TAC-AESOP Task.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Task (project management); Computer science; Identification (biology); Test (biology); Natural language processing; Artificial intelligence; Machine learning; Information retrieval","score_opus":0.02800465364585585,"score_gpt":0.24982617316970912,"score_spread":0.22182151952385326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2403788938","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13308191,0.0007944442,0.8214911,0.00045415547,0.0002768495,0.00076862826,0.0019232426,0.0325536,0.0086560575],"genre_scores_gemma":[0.57776445,0.00021703666,0.40064278,0.0002565148,0.0002443136,0.0007489369,0.006466952,0.0007563695,0.012902699],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99898404,0.00030090127,0.00007503001,0.00037530443,0.00018905332,0.00007559612],"domain_scores_gemma":[0.9981567,0.0006153483,0.00015370428,0.00033248126,0.00063828577,0.00010342646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011205903,0.0007450731,0.0007126562,0.0011178438,0.000588879,0.0011639816,0.0008756054,0.00096970156,0.003751995],"category_scores_gemma":[0.0043312367,0.00023314763,0.0005581612,0.0008114641,0.00020103368,0.0018647237,0.0008964066,0.00060661114,0.004379265],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008005555,0.00065566617,0.006030222,0.00051699707,0.00018958919,0.00031427018,0.00042802023,0.012746558,0.10635474,0.0016043009,0.023736292,0.8466228],"study_design_scores_gemma":[0.00012669052,0.0008654482,0.012158946,0.000056272547,0.00022213158,0.0005926715,0.0004472245,0.8786929,0.08236483,0.005059767,0.019311309,0.00010178462],"about_ca_topic_score_codex":0.0023860263,"about_ca_topic_score_gemma":0.0034627488,"teacher_disagreement_score":0.003751995,"about_ca_system_score_codex":0.00037209422,"about_ca_system_score_gemma":0.0010351918,"threshold_uncertainty_score":0.012551665},"labels":[],"label_agreement":null},{"id":"W2403930574","doi":"","title":"Intelius-NYU TAC-KBP2012 Cold Start System","year":2012,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Task (project management); Computer science; Focus (optics); Cold war; Adaptation (eye); Cold start (automotive); Voting; Start up; Engineering; Politics","score_opus":0.013611818408705267,"score_gpt":0.23897313072513676,"score_spread":0.2253613123164315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2403930574","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013302779,0.0007553486,0.3371593,0.0012677159,0.0007556804,0.001092185,0.12118176,0.4540431,0.07044216],"genre_scores_gemma":[0.06500169,0.00037040305,0.29393435,0.0006900763,0.00028401357,0.0023956462,0.5662997,0.029402046,0.04162214],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99658185,0.0006553821,0.0004584892,0.0007652893,0.0012306081,0.00030844306],"domain_scores_gemma":[0.991663,0.0017977075,0.00047790896,0.0027399715,0.0030240128,0.0002973923],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004114583,0.002116959,0.0019384052,0.004330418,0.0024412654,0.004041886,0.004175608,0.0017794508,0.0512375],"category_scores_gemma":[0.015048175,0.0014158821,0.0010693662,0.00377776,0.00069115084,0.008231215,0.0047626756,0.003051008,0.06941581],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010066209,0.00019244922,0.001741308,0.00067778636,0.00007014881,0.00045127838,0.0008385728,0.0024065734,0.008613944,0.010671026,0.8520848,0.12124539],"study_design_scores_gemma":[0.0003307165,0.00012219555,0.0027698965,0.00021372584,0.00009848989,0.0004520841,0.00058160437,0.09137283,0.028274015,0.012765971,0.862831,0.00018751156],"about_ca_topic_score_codex":0.019621298,"about_ca_topic_score_gemma":0.016320871,"teacher_disagreement_score":0.0512375,"about_ca_system_score_codex":0.0017975972,"about_ca_system_score_gemma":0.0034590191,"threshold_uncertainty_score":0.17140651},"labels":[],"label_agreement":null},{"id":"W2403963403","doi":"","title":"University of Waterloo at the TREC 2013 Temporal Summarization Track.","year":2013,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Cosine similarity; Ranking (information retrieval); Relevance (law); Information retrieval; Similarity (geometry); Metric (unit); Latency (audio); Set (abstract data type); Task (project management); Relevance feedback; Artificial intelligence; Natural language processing; Data mining; Pattern recognition (psychology); Image retrieval","score_opus":0.027573731347673866,"score_gpt":0.21319151961247845,"score_spread":0.18561778826480457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2403963403","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03985818,0.01558799,0.10315083,0.03016127,0.008888322,0.0063272514,0.44194013,0.046640996,0.30744502],"genre_scores_gemma":[0.063940026,0.0038410083,0.12323491,0.0023358364,0.0009883603,0.001975217,0.47793424,0.0041650534,0.32158533],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9951149,0.0009784334,0.00020704263,0.0009220341,0.0023391545,0.00043833177],"domain_scores_gemma":[0.98774683,0.0012107833,0.00024280834,0.0009042881,0.008482391,0.0014128841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069671106,0.0015173727,0.001856497,0.0034285872,0.0030754341,0.004324626,0.0018556319,0.0010892142,0.057038125],"category_scores_gemma":[0.009219456,0.0006532292,0.00062054145,0.0045559816,0.00080117077,0.0031676574,0.0013888745,0.0017285576,0.020410767],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001906149,0.00013841936,0.0004905722,0.0003131988,0.000031581287,0.000047424186,0.00021086694,0.00062149536,0.0044311327,0.0014582963,0.9300441,0.062022284],"study_design_scores_gemma":[0.00025821905,0.00022907373,0.0053833723,0.00011776207,0.000059026966,0.0000548925,0.0004720347,0.011138193,0.012374453,0.0023820035,0.96742445,0.000106390085],"about_ca_topic_score_codex":0.45764893,"about_ca_topic_score_gemma":0.6083322,"teacher_disagreement_score":0.45764893,"about_ca_system_score_codex":0.009968801,"about_ca_system_score_gemma":0.014461805,"threshold_uncertainty_score":0.9099702},"labels":[],"label_agreement":null},{"id":"W2405071382","doi":"10.5281/zenodo.43767","title":"Novel Topic N-Gram Count Lm Incorporating Document-Based Topic Distributions And N-Gram Counts","year":2014,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut National de la Recherche Scientifique","funders":"","keywords":"Gram; n-gram; Statistics; Computer science; Mathematics; Natural language processing; Biology; Bacteria","score_opus":0.027033037076250207,"score_gpt":0.23966099022278342,"score_spread":0.2126279531465332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405071382","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009761942,0.001509163,0.9756797,0.0009580541,0.00083788956,0.00013680336,0.0012738495,0.007575063,0.0022675712],"genre_scores_gemma":[0.140901,0.0012818265,0.82103556,0.0011106443,0.0016621941,0.0010024155,0.0096169375,0.00256102,0.020828513],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976546,0.0010442429,0.00016576979,0.0005387329,0.0004407436,0.00015591533],"domain_scores_gemma":[0.9952857,0.0030376357,0.00014068688,0.00045056528,0.000898745,0.0001867004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003098312,0.0020460237,0.0019730914,0.0014199723,0.001387882,0.00220018,0.0025874549,0.003362433,0.0077236486],"category_scores_gemma":[0.009796181,0.0008839273,0.0013616009,0.0020075715,0.0005201421,0.0034346143,0.0020633168,0.0033681665,0.0093532745],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012310573,0.0003365419,0.0012724856,0.00059894926,0.00043651808,0.00035514295,0.00032666014,0.048470333,0.033679094,0.012588727,0.05116425,0.8495403],"study_design_scores_gemma":[0.00007172823,0.00009130774,0.0005330921,0.000039361217,0.00008942094,0.00013042672,0.00005399079,0.9748317,0.0076167397,0.007365225,0.009131998,0.000045052868],"about_ca_topic_score_codex":0.0052988483,"about_ca_topic_score_gemma":0.012265069,"teacher_disagreement_score":0.0077236486,"about_ca_system_score_codex":0.0009291136,"about_ca_system_score_gemma":0.002546558,"threshold_uncertainty_score":0.025838196},"labels":[],"label_agreement":null},{"id":"W2405155347","doi":"","title":"Getting Emotional About News.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Feeling; Interpretation (philosophy); Computer science; Process (computing); Psychology; Social psychology; Information retrieval","score_opus":0.018652736969281478,"score_gpt":0.2403762818665267,"score_spread":0.22172354489724522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405155347","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14038216,0.023904564,0.15918465,0.049872007,0.006519581,0.0009729251,0.019170094,0.007500229,0.5924938],"genre_scores_gemma":[0.65600455,0.013503886,0.079676256,0.008134194,0.0024841295,0.00033802364,0.016243529,0.0012930762,0.22232229],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986928,0.00040979186,0.000086141794,0.00024140916,0.00048228243,0.00008753597],"domain_scores_gemma":[0.99479645,0.0021811444,0.0006323732,0.0005839704,0.0013583131,0.00044782934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020366644,0.00048055753,0.00025528934,0.0016831909,0.0011894322,0.0040416406,0.0003610959,0.0012176196,0.021452833],"category_scores_gemma":[0.019700697,0.00020389359,0.00031047137,0.0014105794,0.00067290413,0.0045367517,0.0015756973,0.0011060634,0.012776563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023560206,0.000073797804,0.015732454,0.0010445908,0.0001038188,0.0002780827,0.020997034,0.00045890032,0.007385029,0.035767406,0.30098063,0.61694264],"study_design_scores_gemma":[0.000014571662,0.0001632692,0.044004224,0.0004360269,0.00013007871,0.0012680065,0.01803655,0.0029232637,0.0043878695,0.031249434,0.8972776,0.00010903127],"about_ca_topic_score_codex":0.0024217062,"about_ca_topic_score_gemma":0.006546124,"teacher_disagreement_score":0.021452833,"about_ca_system_score_codex":0.0007589937,"about_ca_system_score_gemma":0.00044403187,"threshold_uncertainty_score":0.07176691},"labels":[],"label_agreement":null},{"id":"W2405247712","doi":"","title":"TJU GSummary at TAC2011: Category oriented extractive content selection for guided summarization.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Selection (genetic algorithm); Computer science; Content (measure theory); Mathematics; Natural language processing; Artificial intelligence","score_opus":0.04326973547035373,"score_gpt":0.2589672449241649,"score_spread":0.21569750945381116,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405247712","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01666352,0.0031545267,0.88559705,0.005926796,0.0035603854,0.00090853905,0.01585345,0.054193217,0.014142577],"genre_scores_gemma":[0.09748995,0.0015087412,0.793397,0.0013605775,0.0014590275,0.0011401975,0.043436736,0.011013996,0.049193736],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99670655,0.0013923233,0.0001798553,0.0005805393,0.0009194422,0.00022137289],"domain_scores_gemma":[0.9915386,0.0034319847,0.00018517081,0.0011661344,0.0030850207,0.0005930667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004357969,0.0017753873,0.001729445,0.004128126,0.0014984973,0.0034417456,0.0017333893,0.0027018432,0.021749003],"category_scores_gemma":[0.0156135345,0.00063878705,0.0011257356,0.003691529,0.00065865536,0.0043369015,0.0025435814,0.0023301134,0.019448148],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010363593,0.00030042083,0.0006940273,0.0007638404,0.00016640444,0.00046133308,0.00094826694,0.0034517783,0.047064513,0.0074432595,0.4766313,0.46103856],"study_design_scores_gemma":[0.00060720986,0.00072222645,0.0044114245,0.0002900823,0.0003870743,0.00074418867,0.0012871874,0.23965557,0.13679072,0.047051374,0.5676669,0.00038604432],"about_ca_topic_score_codex":0.004314553,"about_ca_topic_score_gemma":0.0069263675,"teacher_disagreement_score":0.021749003,"about_ca_system_score_codex":0.0007646114,"about_ca_system_score_gemma":0.001594596,"threshold_uncertainty_score":0.0727576},"labels":[],"label_agreement":null},{"id":"W2406782120","doi":"","title":"An Accuracy-Oriented Divide-and-Conquer Strategy for Recognizing Textual Entailment.","year":2008,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Divide and conquer algorithms; Physics; Combinatorics; Mathematics; Algorithm","score_opus":0.03076893746456134,"score_gpt":0.2896070888439543,"score_spread":0.258838151379393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2406782120","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026428998,0.0007581887,0.93332005,0.0014744749,0.00020547934,0.0008914285,0.0023094008,0.019600175,0.0150117725],"genre_scores_gemma":[0.13524769,0.00016611593,0.8441292,0.00043465532,0.00014586328,0.00041439,0.005994632,0.001370939,0.012096539],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958281,0.0009864885,0.00051103556,0.001168354,0.0011286038,0.00037742505],"domain_scores_gemma":[0.99489236,0.0021411872,0.00026497032,0.0012087565,0.0012631036,0.0002295976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003410224,0.0026301825,0.0015963233,0.0033050822,0.001539161,0.0027320592,0.0033264584,0.0025632153,0.023593746],"category_scores_gemma":[0.01666118,0.00073350786,0.0019010648,0.0027018753,0.001021242,0.004301297,0.0036711884,0.0027107901,0.011887501],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085969473,0.0003729091,0.0018720619,0.00037809528,0.00016762491,0.0001878843,0.0003626023,0.0055996305,0.019227264,0.010618588,0.04688864,0.913465],"study_design_scores_gemma":[0.00038062144,0.00058850186,0.0039940057,0.00014499333,0.00042442686,0.00094313297,0.0013404187,0.75895596,0.075918846,0.10395414,0.05323635,0.00011865358],"about_ca_topic_score_codex":0.0057753436,"about_ca_topic_score_gemma":0.012992766,"teacher_disagreement_score":0.023593746,"about_ca_system_score_codex":0.0014820413,"about_ca_system_score_gemma":0.002926448,"threshold_uncertainty_score":0.07892895},"labels":[],"label_agreement":null},{"id":"W2408078294","doi":"","title":"Using a weakly supervised approach and lexical patterns for the KBP slot filling task.","year":2011,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Task (project management); Computer science; Identification (biology); Relation (database); Artificial intelligence; Natural language processing; Population; Track (disk drive); Knowledge base; Component (thermodynamics); Data mining; Engineering; Biology","score_opus":0.06256469848337742,"score_gpt":0.267818551467447,"score_spread":0.20525385298406956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408078294","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028708162,0.00010035558,0.95924234,0.00025359195,0.000043581145,0.0004906285,0.0008234994,0.006759286,0.003578471],"genre_scores_gemma":[0.17181742,0.000113525915,0.8133335,0.0002090448,0.000072082905,0.0007813572,0.0057823607,0.0006112526,0.007279484],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99703884,0.0011315866,0.00020989098,0.00087259297,0.0006364692,0.000110599816],"domain_scores_gemma":[0.99261606,0.0046115955,0.00046944834,0.0011646221,0.0009190419,0.00021928028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024758442,0.0009951825,0.0009117016,0.0022794378,0.0009874433,0.002566454,0.0017838452,0.0015244756,0.0034629528],"category_scores_gemma":[0.015709503,0.0005527336,0.00097172055,0.0017264327,0.00062906364,0.0043159993,0.00217209,0.0014175734,0.005607406],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006962992,0.0009246779,0.011048802,0.00064335426,0.00028667267,0.00037879177,0.002215272,0.010549269,0.093208656,0.009999464,0.018405266,0.85164344],"study_design_scores_gemma":[0.00015799496,0.0005264241,0.011610184,0.00016626863,0.0002964494,0.0009433684,0.0016355552,0.78349286,0.08173982,0.06613977,0.053127114,0.00016422562],"about_ca_topic_score_codex":0.002063259,"about_ca_topic_score_gemma":0.0045265793,"teacher_disagreement_score":0.0034629528,"about_ca_system_score_codex":0.00042531276,"about_ca_system_score_gemma":0.0015867082,"threshold_uncertainty_score":0.01309365},"labels":[],"label_agreement":null},{"id":"W2408135015","doi":"10.1016/j.compedu.2016.05.007","title":"Automatic detection of expert models: The exploration of expert modeling methods applicable to technology-based assessment and instruction","year":2016,"lang":"en","type":"article","venue":"Computers & Education","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Expert system; Artificial intelligence; Machine learning; Data science","score_opus":0.06050415620251584,"score_gpt":0.3599593436345837,"score_spread":0.2994551874320679,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408135015","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022542326,0.00030480433,0.9748584,0.0002121195,0.000017633212,0.000067001085,0.00012733649,0.0009307447,0.0009395819],"genre_scores_gemma":[0.43242887,0.00035877852,0.5645249,0.00010842551,0.000044441585,0.00016471853,0.0004688592,0.00024426516,0.0016568264],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99747676,0.0012179028,0.00011983104,0.00045830238,0.00060563255,0.00012158933],"domain_scores_gemma":[0.981138,0.014663151,0.00082071783,0.0013931175,0.0017236763,0.00026134198],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039744647,0.0008014815,0.0009766942,0.0018891251,0.000447942,0.0024366167,0.0018547606,0.0013425403,0.0016717982],"category_scores_gemma":[0.024649309,0.00048732336,0.0008378852,0.0009916518,0.00049268856,0.0025430492,0.0014467776,0.0015856178,0.0007075217],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030954485,0.0004620398,0.010567884,0.0004401681,0.00018671852,0.00014895487,0.000934918,0.08384872,0.017825117,0.02005684,0.003973301,0.8612459],"study_design_scores_gemma":[0.00001813656,0.000053069227,0.0024511772,0.00003816328,0.000033698703,0.00012047088,0.00013428276,0.9640476,0.0064284634,0.024881482,0.0017688917,0.000024620338],"about_ca_topic_score_codex":0.0031415005,"about_ca_topic_score_gemma":0.0041376622,"teacher_disagreement_score":0.0039744647,"about_ca_system_score_codex":0.00059712277,"about_ca_system_score_gemma":0.0014464443,"threshold_uncertainty_score":0.02101922},"labels":[],"label_agreement":null},{"id":"W2408142165","doi":"","title":"Perspectives on Ochre Provenance in British Columbia, Canada","year":2016,"lang":"en","type":"article","venue":"The 81st Annual Meeting of the Society for American Archaeology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Provenance; Archaeology; History; Geology; Geochemistry","score_opus":0.005980302456173531,"score_gpt":0.2140235990248641,"score_spread":0.20804329656869056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408142165","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13853991,0.023106962,0.016806569,0.2918307,0.002822874,0.00027650496,0.004502939,0.0008541586,0.52125937],"genre_scores_gemma":[0.7463589,0.016689222,0.0140958,0.011090386,0.0005285246,0.00011021487,0.0018918727,0.00069752935,0.20853767],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9938539,0.0011592836,0.00045721326,0.0006998572,0.00226903,0.0015607305],"domain_scores_gemma":[0.9600907,0.010281768,0.0011297241,0.0022792257,0.022960812,0.0032577456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009993863,0.00040976293,0.0008048843,0.0052634506,0.022317434,0.020051269,0.0025077392,0.0023354492,0.013955264],"category_scores_gemma":[0.031065082,0.0007070933,0.0004155475,0.0111195585,0.00984629,0.0064817746,0.0041325414,0.0038277113,0.0008806171],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020580315,0.000060790524,0.016258381,0.0003962532,0.000041983898,0.0013111664,0.06425151,0.0031168354,0.0010088703,0.52554363,0.18785377,0.19995095],"study_design_scores_gemma":[0.000026571131,0.000015340107,0.014506434,0.0007146983,0.000035724315,0.0001964697,0.04323622,0.0019425615,0.00052799407,0.045669228,0.89297223,0.00015656446],"about_ca_topic_score_codex":0.99708754,"about_ca_topic_score_gemma":0.9985072,"teacher_disagreement_score":0.14097941,"about_ca_system_score_codex":0.14097941,"about_ca_system_score_gemma":0.26330975,"threshold_uncertainty_score":0.9963421},"labels":[],"label_agreement":null},{"id":"W2408372179","doi":"","title":"Zero-Shot Learning and Clustering for Semantic Utterance Classification","year":2013,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Cluster analysis; Utterance; Discriminative model; Natural language processing; Task (project management); Pattern recognition (psychology); Machine learning","score_opus":0.09195313340748311,"score_gpt":0.19559357944619718,"score_spread":0.10364044603871407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408372179","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02068023,0.0003736893,0.9757888,0.0001948476,0.00007941431,0.000095234485,0.00020580637,0.001378822,0.0012032037],"genre_scores_gemma":[0.51662534,0.00029325808,0.47543865,0.00031313262,0.00020939931,0.00031605086,0.0023048155,0.00026507638,0.004234325],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980692,0.0005552048,0.00010008961,0.0006090389,0.00043810462,0.00022844152],"domain_scores_gemma":[0.9983902,0.00066452334,0.00011586233,0.00033137656,0.00038143212,0.000116712035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013229708,0.0012353708,0.0015911206,0.0018198052,0.0008543566,0.0011456121,0.0031039594,0.0016171403,0.0026822686],"category_scores_gemma":[0.0044922405,0.00044798345,0.0010285376,0.0014062498,0.0011993936,0.002617146,0.0020826012,0.0023098225,0.0012637663],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064665516,0.0006724971,0.003073924,0.00033553163,0.00017087121,0.00014547181,0.0006361047,0.057233043,0.019348588,0.023761623,0.011749566,0.88222617],"study_design_scores_gemma":[0.00002076632,0.00011786551,0.00080439186,0.000018913757,0.00002538734,0.000074826334,0.00012965492,0.96773046,0.008543546,0.02046481,0.0020339396,0.00003538391],"about_ca_topic_score_codex":0.0054735634,"about_ca_topic_score_gemma":0.006864054,"teacher_disagreement_score":0.0054735634,"about_ca_system_score_codex":0.0010780406,"about_ca_system_score_gemma":0.001321084,"threshold_uncertainty_score":0.010883391},"labels":[],"label_agreement":null},{"id":"W2409591106","doi":"10.18653/v1/d16-1147","title":"Key-Value Memory Networks for Directly Reading Documents","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":204,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Schema (genetic algorithms); Key (lock); Reading (process); Information retrieval; Construct (python library); Benchmark (surveying); Domain (mathematical analysis); Value (mathematics); Artificial intelligence; Machine learning; Programming language","score_opus":0.025207387836240547,"score_gpt":0.28309490330139714,"score_spread":0.2578875154651566,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2409591106","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039406992,0.002412638,0.9330784,0.0013667188,0.00015250857,0.00023986169,0.004185383,0.010636889,0.008520595],"genre_scores_gemma":[0.33772573,0.001437211,0.6338562,0.00041660713,0.00022627847,0.0006390142,0.011017378,0.0012661824,0.013415401],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99894625,0.0003152833,0.00008433215,0.00036707366,0.00020192201,0.00008517606],"domain_scores_gemma":[0.99410653,0.0037844854,0.0003732372,0.0011516981,0.00048952445,0.000094471085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013295253,0.0014614669,0.0008512301,0.0025338312,0.00091000495,0.0023825136,0.0025662233,0.0019390187,0.007613529],"category_scores_gemma":[0.012645449,0.0006687683,0.00085681217,0.0027848948,0.00068427896,0.011079102,0.0019791347,0.0018952561,0.0037866335],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041891233,0.00021842889,0.0027652355,0.00086809933,0.00012284247,0.00022222481,0.0009260525,0.07884081,0.011756651,0.0646362,0.045294754,0.7939299],"study_design_scores_gemma":[0.000054291075,0.0000898646,0.00072353997,0.00010464748,0.00008432906,0.00024944782,0.0004206053,0.72693574,0.021414397,0.21640112,0.03347404,0.00004797892],"about_ca_topic_score_codex":0.0038851039,"about_ca_topic_score_gemma":0.0070774006,"teacher_disagreement_score":0.007613529,"about_ca_system_score_codex":0.0015393127,"about_ca_system_score_gemma":0.0011446625,"threshold_uncertainty_score":0.02546984},"labels":[],"label_agreement":null},{"id":"W2418993857","doi":"10.1609/aaai.v31i1.10984","title":"Multiresolution Recurrent Neural Networks: An Application to Dialogue Response Generation","year":2017,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":162,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Margin (machine learning); Recurrent neural network; Sequence (biology); Artificial intelligence; Semantics (computer science); Representation (politics); Artificial neural network; Abstraction; Machine learning; Process (computing); Natural language processing","score_opus":0.13868579761646666,"score_gpt":0.3393223618512201,"score_spread":0.20063656423475346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2418993857","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024820482,0.00059214956,0.969738,0.00032154232,0.00006074532,0.00006486763,0.00014831717,0.0024103338,0.0018435791],"genre_scores_gemma":[0.650272,0.0005071396,0.34412864,0.00024406402,0.000088080174,0.00021066237,0.00042069177,0.0002756844,0.003852978],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99938416,0.00027613487,0.000028344408,0.00014174827,0.00012313634,0.00004637521],"domain_scores_gemma":[0.9988686,0.0007271959,0.00010794133,0.00011962804,0.00013539108,0.000041204188],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013881682,0.0008621636,0.0006533443,0.00050310965,0.00023230512,0.00072367315,0.0012641462,0.001001679,0.0019930245],"category_scores_gemma":[0.004518589,0.00044242907,0.0007739615,0.0005547595,0.00037588354,0.0012665711,0.000815781,0.0013736122,0.0006132079],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023683021,0.00014349582,0.0009158442,0.00016457276,0.00011464918,0.0002728406,0.00032741218,0.77493894,0.016646605,0.015058014,0.0028661585,0.18831456],"study_design_scores_gemma":[0.0000041057365,0.00001647762,0.000054629876,0.0000033447177,0.0000064125284,0.000015158672,0.0000041653166,0.9965417,0.0009340644,0.0020455876,0.00036997735,0.000004391177],"about_ca_topic_score_codex":0.0040321257,"about_ca_topic_score_gemma":0.0045395587,"teacher_disagreement_score":0.0040321257,"about_ca_system_score_codex":0.0006753828,"about_ca_system_score_gemma":0.00039842873,"threshold_uncertainty_score":0.008017302},"labels":[],"label_agreement":null},{"id":"W2431165553","doi":"10.1145/1645953.1646274","title":"Answer typing for information retrieval","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Discriminative model; Computer science; Information retrieval; Ranking (information retrieval); Typing; Classifier (UML); Selection (genetic algorithm); Identification (biology); Class (philosophy); Artificial intelligence; Natural language processing","score_opus":0.019855429461452807,"score_gpt":0.25267395003336574,"score_spread":0.23281852057191293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2431165553","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003721602,0.01302963,0.965175,0.0034142158,0.00045522177,0.0003021107,0.0008985957,0.0035792857,0.009424426],"genre_scores_gemma":[0.16563216,0.010278624,0.80554837,0.002234285,0.0018213718,0.0007406038,0.002955417,0.00065256155,0.010136668],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9917209,0.003648441,0.0007193034,0.0014551098,0.0021687082,0.00028765382],"domain_scores_gemma":[0.9863788,0.007262887,0.0008960513,0.0038971198,0.0012917796,0.0002734464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068376455,0.0013828625,0.0018772188,0.0045996215,0.0011392528,0.004980098,0.0022113475,0.0028560734,0.0096255895],"category_scores_gemma":[0.031191235,0.000682286,0.001746249,0.0060073384,0.00244366,0.009398045,0.002587081,0.0026766441,0.009086179],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019315213,0.00015255713,0.0017161138,0.0013081216,0.00011882067,0.00017187532,0.00051869935,0.012289038,0.004613126,0.31627822,0.039179385,0.6234609],"study_design_scores_gemma":[0.000036137397,0.00012128808,0.0010018576,0.00033416148,0.000093496245,0.00060849584,0.0001957355,0.18636975,0.0033511783,0.7127563,0.095034316,0.000097302065],"about_ca_topic_score_codex":0.0023761548,"about_ca_topic_score_gemma":0.0015908782,"teacher_disagreement_score":0.0096255895,"about_ca_system_score_codex":0.0018298867,"about_ca_system_score_gemma":0.0014347723,"threshold_uncertainty_score":0.036161363},"labels":[],"label_agreement":null},{"id":"W2464868641","doi":"10.1145/2911451.2917767","title":"SIGIR 2016 Workshop WebQA II","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science","score_opus":0.02454648631855706,"score_gpt":0.24327230933541028,"score_spread":0.21872582301685323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2464868641","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020814335,0.12206046,0.37506652,0.107272156,0.109052196,0.0034492982,0.017243212,0.035204373,0.20983748],"genre_scores_gemma":[0.07983667,0.035790678,0.1771132,0.024406293,0.025221493,0.002516964,0.078688726,0.011177607,0.56524837],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98759913,0.005118988,0.0009246043,0.0019875397,0.0035783716,0.0007913429],"domain_scores_gemma":[0.9869421,0.0029085858,0.0002834841,0.002249779,0.00571153,0.0019043853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022703545,0.0028153062,0.003411219,0.0036421588,0.0028931112,0.009908013,0.00384388,0.004882492,0.070731506],"category_scores_gemma":[0.021167248,0.0012582307,0.0025963713,0.0028758405,0.0019292469,0.01124067,0.0057995757,0.006515239,0.092905916],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003453182,0.00024880285,0.00033684602,0.000494897,0.00008504838,0.00015855388,0.00038342096,0.0009487764,0.0033206441,0.0062896563,0.8029591,0.18442902],"study_design_scores_gemma":[0.000121732766,0.00015593397,0.0008037812,0.00038395452,0.000051225783,0.00023918643,0.00058825413,0.007729067,0.0029711078,0.019786188,0.9671107,0.000058969734],"about_ca_topic_score_codex":0.012191686,"about_ca_topic_score_gemma":0.010718617,"teacher_disagreement_score":0.070731506,"about_ca_system_score_codex":0.003697937,"about_ca_system_score_gemma":0.0069202003,"threshold_uncertainty_score":0.23662049},"labels":[],"label_agreement":null},{"id":"W2465765001","doi":"10.1145/2964797.2964806","title":"Report on the Eighth Workshop on Exploiting Semantic Annotations in Information Retrieval (ESAIR '15)","year":2016,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada)","funders":"","keywords":"Computer science; Event (particle physics); Focus (optics); World Wide Web; Field (mathematics); Semantic Web; Track (disk drive); Data science; Information retrieval","score_opus":0.030702390798121124,"score_gpt":0.262404220364391,"score_spread":0.23170182956626986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2465765001","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033188697,0.04627971,0.39953712,0.11104226,0.11749884,0.0068038567,0.048829053,0.021433014,0.21538745],"genre_scores_gemma":[0.051096015,0.014058311,0.21950309,0.01710662,0.014009937,0.0034170682,0.12378228,0.0069330707,0.5500936],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98840415,0.004846534,0.0005345815,0.0016220679,0.0034410052,0.0011516047],"domain_scores_gemma":[0.9759261,0.0067657237,0.0004540107,0.0036100156,0.008171106,0.0050730347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023179872,0.0021987653,0.002259471,0.0028970675,0.0022853187,0.0089134425,0.0037239122,0.00404008,0.09642976],"category_scores_gemma":[0.024921842,0.0007564941,0.002072697,0.0028116463,0.0012175918,0.012884036,0.008711067,0.005138512,0.065905906],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037527492,0.00047413996,0.00045618278,0.00040796393,0.000060979357,0.00009122564,0.00065415655,0.0004964106,0.0031940693,0.0033627944,0.87855273,0.11187406],"study_design_scores_gemma":[0.00013593867,0.00024651468,0.0014514445,0.0002662139,0.00007817983,0.000121769845,0.0008477791,0.0031581502,0.0033170613,0.007671389,0.9826302,0.00007546252],"about_ca_topic_score_codex":0.009639099,"about_ca_topic_score_gemma":0.014722726,"teacher_disagreement_score":0.09642976,"about_ca_system_score_codex":0.002461338,"about_ca_system_score_gemma":0.004922345,"threshold_uncertainty_score":0.3225897},"labels":[],"label_agreement":null},{"id":"W2468158014","doi":"10.18653/v1/s16-1133","title":"Overfitting at SemEval-2016 Task 3: Detecting Semantically Similar Questions in Community Question Answering Forums with Word Embeddings","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"SemEval; Computer science; Question answering; Artificial intelligence; Similarity (geometry); Task (project management); Overfitting; Word (group theory); Natural language processing; Support vector machine; Set (abstract data type); WordNet; F1 score; Semantic similarity; Binary classification; Test set; Artificial neural network; Word embedding; Test (biology); Embedding; Mathematics","score_opus":0.013729434556114037,"score_gpt":0.2543705769289579,"score_spread":0.2406411423728439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2468158014","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6890742,0.005874057,0.26926836,0.0013644149,0.0012147174,0.0008304709,0.007420668,0.019782854,0.005170351],"genre_scores_gemma":[0.87170166,0.00026264202,0.09645393,0.00048388782,0.0002837417,0.00042435495,0.025866369,0.0005462895,0.003977122],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99022746,0.00495253,0.00062231626,0.0025142496,0.0011284908,0.0005549146],"domain_scores_gemma":[0.9857525,0.0085579315,0.00073510857,0.002494065,0.0019300986,0.00053026364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012493922,0.0032110873,0.0021522995,0.0030339255,0.0013364874,0.0020089056,0.002152907,0.0036286658,0.002421752],"category_scores_gemma":[0.031885516,0.00040335231,0.0022734508,0.0017667431,0.0008008959,0.0040708585,0.0033042873,0.002440894,0.0026291758],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035113364,0.0026393265,0.08740182,0.002008679,0.0015180295,0.0011722646,0.0021130505,0.06139084,0.043561053,0.003150427,0.07476473,0.71676844],"study_design_scores_gemma":[0.00036774148,0.0013296872,0.03845762,0.00021646952,0.00037498507,0.0017190817,0.0015982388,0.8782904,0.037909783,0.015448609,0.02411866,0.00016864322],"about_ca_topic_score_codex":0.004558493,"about_ca_topic_score_gemma":0.0076559605,"teacher_disagreement_score":0.012493922,"about_ca_system_score_codex":0.0008429118,"about_ca_system_score_gemma":0.0010349364,"threshold_uncertainty_score":0.06607497},"labels":[],"label_agreement":null},{"id":"W2468590752","doi":"10.18653/v1/s16-1118","title":"DalGTM at SemEval-2016 Task 1: Importance-Aware Compositional Approach to Short Text Similarity","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada; Boeing","keywords":"SemEval; Computer science; Similarity (geometry); Word (group theory); Semantic similarity; Pairwise comparison; Natural language processing; Context (archaeology); Task (project management); Artificial intelligence; Rank (graph theory); Information retrieval; Image (mathematics); Mathematics","score_opus":0.031227635079748378,"score_gpt":0.2505197466858442,"score_spread":0.21929211160609582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2468590752","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23762816,0.0062998943,0.36220372,0.003326999,0.008149443,0.0064748544,0.067556635,0.269838,0.038522188],"genre_scores_gemma":[0.3641075,0.0007010647,0.40720177,0.00089635054,0.0009982447,0.0030730427,0.18170965,0.010212046,0.03110041],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912224,0.0033184153,0.00075627834,0.0023443128,0.0018787091,0.00047996026],"domain_scores_gemma":[0.99207723,0.0019386903,0.0002842794,0.0019137189,0.0027879851,0.0009981407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006982636,0.0029718347,0.0029267513,0.0035266494,0.0019306752,0.003639402,0.0034201741,0.00319589,0.01760171],"category_scores_gemma":[0.024868205,0.00067748374,0.0018049034,0.0019483791,0.0007323142,0.0052449433,0.00588816,0.0026697717,0.01964787],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032336195,0.0015122873,0.0048129237,0.0025661618,0.00060096575,0.00092042953,0.0017646253,0.009593137,0.04214521,0.004696369,0.3675034,0.5606508],"study_design_scores_gemma":[0.0015247382,0.0030932287,0.019319097,0.00027822412,0.00036089311,0.0025860942,0.0026376562,0.42513013,0.09324358,0.025307698,0.4259563,0.00056246575],"about_ca_topic_score_codex":0.005494095,"about_ca_topic_score_gemma":0.0074460856,"teacher_disagreement_score":0.01760171,"about_ca_system_score_codex":0.0016155466,"about_ca_system_score_gemma":0.0023358124,"threshold_uncertainty_score":0.058883548},"labels":[],"label_agreement":null},{"id":"W2469016277","doi":"10.1007/978-3-319-34111-8_20","title":"Harnessing Open Information Extraction for Entity Classification in a French Corpus","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Open domain; Information extraction; Domain (mathematical analysis); Precision and recall; Information retrieval; Natural language processing; Recall; Named-entity recognition; Relationship extraction; Artificial intelligence; Entity linking; Question answering; Knowledge base; Task (project management); Linguistics","score_opus":0.0454656989705401,"score_gpt":0.2937363488754873,"score_spread":0.24827064990494718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2469016277","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21824592,0.015650274,0.6094315,0.006291649,0.002488177,0.0007997962,0.07285876,0.029397994,0.04483592],"genre_scores_gemma":[0.3881439,0.0038553807,0.39889368,0.00064665,0.0010487916,0.00074874057,0.18590088,0.0031928231,0.01756915],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99702877,0.0011009604,0.00029909532,0.0008675885,0.000491617,0.00021209741],"domain_scores_gemma":[0.9893923,0.006673234,0.00029879078,0.0012478316,0.0021887484,0.00019920105],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032216252,0.0014866806,0.0010216973,0.007031884,0.0026002936,0.0038797394,0.0011216865,0.0012729306,0.0068972926],"category_scores_gemma":[0.011493762,0.00069820444,0.0011936818,0.005052744,0.0010463938,0.005197991,0.0029731393,0.0017511952,0.0048903343],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056780345,0.00027731803,0.007654846,0.00179539,0.00024023735,0.0017623708,0.004158619,0.008143332,0.040541194,0.022724105,0.1096199,0.802515],"study_design_scores_gemma":[0.00020182309,0.0002889312,0.034855157,0.00076290464,0.00068693276,0.00255461,0.004620643,0.16288094,0.076249264,0.040154193,0.6763362,0.0004083141],"about_ca_topic_score_codex":0.03338187,"about_ca_topic_score_gemma":0.041209307,"teacher_disagreement_score":0.03338187,"about_ca_system_score_codex":0.0018838051,"about_ca_system_score_gemma":0.0030968692,"threshold_uncertainty_score":0.06637514},"labels":[],"label_agreement":null},{"id":"W2469060249","doi":"10.18653/v1/n16-1108","title":"Pairwise Word Interaction Modeling with Deep Neural Networks for Semantic Similarity Measurement","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":249,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Science Foundation","keywords":"Computer science; Artificial intelligence; SemEval; Pairwise comparison; Similarity (geometry); Natural language processing; Semantic similarity; Word (group theory); Sentence; Semantics (computer science); Focus (optics); Artificial neural network; Selection (genetic algorithm); Mathematics","score_opus":0.06594859114948542,"score_gpt":0.25443472665570105,"score_spread":0.18848613550621562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2469060249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051182985,0.00066512194,0.94328564,0.00030034955,0.00007980983,0.000065234126,0.00039164635,0.0018964082,0.0021328265],"genre_scores_gemma":[0.8077451,0.00036561865,0.18590687,0.0002513886,0.00012417775,0.00025021756,0.0014886395,0.00021802043,0.003649857],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991273,0.00025643426,0.000053691187,0.00027694387,0.00020259095,0.00008301179],"domain_scores_gemma":[0.9991449,0.00039975927,0.00012577989,0.0001212125,0.00016292019,0.000045434706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010844738,0.0013110277,0.00090680027,0.0011763861,0.00044177586,0.00095937157,0.0019363933,0.0011790443,0.0025918027],"category_scores_gemma":[0.003792297,0.000336842,0.0007605235,0.0016340215,0.000473515,0.0035062113,0.0015689229,0.002259257,0.00090825907],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053631864,0.00044916625,0.004463384,0.0003572618,0.0004131364,0.00021827148,0.0003972073,0.43416932,0.03605181,0.0351723,0.009083742,0.47868806],"study_design_scores_gemma":[0.0000041758017,0.00001846977,0.00023871052,0.0000034662355,0.000012173475,0.000011949306,0.000011806838,0.98659116,0.0015562274,0.011200286,0.00034599323,0.0000055741248],"about_ca_topic_score_codex":0.0043661427,"about_ca_topic_score_gemma":0.0068767206,"teacher_disagreement_score":0.0043661427,"about_ca_system_score_codex":0.0010661347,"about_ca_system_score_gemma":0.0007343958,"threshold_uncertainty_score":0.0086814165},"labels":[],"label_agreement":null},{"id":"W2473104040","doi":"10.18653/v1/n16-1106","title":"DAG-Structured Long Short-Term Memory for Semantic Compositionality","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; National Research Council Canada","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Principle of compositionality; Computer science; Natural language processing; Semantics (computer science); Artificial intelligence; Representation (politics); Sentence; Semantic memory; Term (time); Recurrent neural network; Sequence (biology); Theoretical computer science; Artificial neural network; Cognition; Programming language","score_opus":0.03229492923303824,"score_gpt":0.28031323163934374,"score_spread":0.2480183024063055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2473104040","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04244975,0.0006201696,0.9516463,0.0005194095,0.00012177568,0.00003886984,0.00050707505,0.0018134359,0.0022831385],"genre_scores_gemma":[0.8585404,0.00081417355,0.1352253,0.00024068027,0.0000763236,0.00009284167,0.0012122385,0.00015097624,0.0036470026],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99977154,0.000066316774,0.000017880318,0.00007986836,0.00003409683,0.000030323314],"domain_scores_gemma":[0.9992999,0.00036330475,0.00008713487,0.0001159701,0.000099614066,0.000034211265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006878302,0.00084628747,0.000574136,0.00076059165,0.00031860493,0.00067266496,0.00107549,0.0006384833,0.0028289198],"category_scores_gemma":[0.0028393217,0.00032449525,0.0007247888,0.00094217015,0.00045933478,0.0027150372,0.0007955764,0.0014780986,0.00073897414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003558298,0.00019764401,0.0026476118,0.00043418948,0.00020387562,0.00028918224,0.00031333842,0.5520325,0.019520463,0.09390593,0.0059700236,0.32412943],"study_design_scores_gemma":[0.000007075674,0.00002657986,0.00021623675,0.000009426889,0.000024570354,0.000027235315,0.000014161089,0.9477665,0.0017980282,0.049026534,0.0010754272,0.000008260563],"about_ca_topic_score_codex":0.004402284,"about_ca_topic_score_gemma":0.007878447,"teacher_disagreement_score":0.004402284,"about_ca_system_score_codex":0.0008114533,"about_ca_system_score_gemma":0.00085376936,"threshold_uncertainty_score":0.0094637275},"labels":[],"label_agreement":null},{"id":"W247476514","doi":"","title":"Absolutist, Pragmatist and Realist Approaches to Research Ethics in the Digital Humanities: The Case of the Schneerson Collection","year":2013,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Pragmatism; Realism; Digital humanities; Judaism; Data collection; Sociology; Art; Epistemology; Philosophy; Social science; Humanities; Theology","score_opus":0.16479174223105567,"score_gpt":0.3114963985367503,"score_spread":0.14670465630569465,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W247476514","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09767382,0.0047427104,0.19220026,0.29008275,0.001058649,0.0008326205,0.00012821106,0.00013316338,0.41314778],"genre_scores_gemma":[0.9275061,0.0007494246,0.037194792,0.012465932,0.0003082236,0.0008604217,0.000025888863,0.00008173931,0.020807464],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.84684193,0.12368171,0.003198977,0.0071658334,0.015394575,0.0037169703],"domain_scores_gemma":[0.9169607,0.056939803,0.0056401216,0.013765432,0.0042574177,0.0024365324],"candidate_categories":["research_integrity"],"consensus_categories":[],"category_scores_codex":[0.098446555,0.00065007724,0.00083176914,0.002550597,0.02113395,0.020103369,0.00221533,0.010832363,0.0025862593],"category_scores_gemma":[0.071560495,0.00088498025,0.0009853852,0.0025176695,0.10969611,0.017031165,0.0121067995,0.012366674,0.0005703546],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011390992,0.000010662442,0.00026455603,0.00003129425,0.0000054279285,0.00014964091,0.022252778,0.00010371622,0.00006110236,0.9735474,0.0013813756,0.002180684],"study_design_scores_gemma":[0.000044222812,0.000037394504,0.00065408496,0.0003157394,0.000016821816,0.00067105045,0.016107025,0.0012330854,0.0005808632,0.8765803,0.10369642,0.000062899104],"about_ca_topic_score_codex":0.009496005,"about_ca_topic_score_gemma":0.010865634,"teacher_disagreement_score":0.98916763,"about_ca_system_score_codex":0.014444878,"about_ca_system_score_gemma":0.013967644,"threshold_uncertainty_score":0.5206413},"labels":[],"label_agreement":null},{"id":"W2476992958","doi":"10.1007/978-3-319-41754-7_46","title":"Automatic Text Summarization with a Reduced Vocabulary Using Continuous Space Vectors","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Agence Nationale de la Recherche","keywords":"Automatic summarization; Computer science; Vocabulary; Natural language processing; Artificial intelligence; Context (archaeology); State (computer science); Space (punctuation); Information retrieval; Speech recognition; Algorithm; Linguistics","score_opus":0.015712199576776567,"score_gpt":0.23180699008876857,"score_spread":0.21609479051199199,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2476992958","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02736607,0.003318204,0.9460103,0.00038886187,0.0009757391,0.00044120947,0.0042079305,0.0143336095,0.0029581983],"genre_scores_gemma":[0.17562804,0.0020699182,0.7780552,0.00019162042,0.0009994456,0.00080583506,0.027982103,0.0012871935,0.012980694],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987984,0.00023890992,0.00015139981,0.00039831252,0.00029150757,0.00012152561],"domain_scores_gemma":[0.99875724,0.0003590031,0.00008525722,0.00018255562,0.00055994827,0.000055954602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070375245,0.0017468426,0.0017168875,0.0031986206,0.00078929076,0.0019438483,0.0011262054,0.0010504149,0.0069109653],"category_scores_gemma":[0.0025341723,0.0004472873,0.0014868305,0.0028670004,0.00035118088,0.002478829,0.00162438,0.0014079242,0.0076882117],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006371153,0.0001399638,0.0003251343,0.000683622,0.00014052486,0.00016480932,0.00021230242,0.004724362,0.11236831,0.0025881159,0.02868491,0.84933084],"study_design_scores_gemma":[0.00033660146,0.0013335616,0.004500234,0.00026888467,0.00090896257,0.0009532448,0.0012652485,0.719545,0.15331732,0.01873046,0.09862638,0.0002140916],"about_ca_topic_score_codex":0.0024561395,"about_ca_topic_score_gemma":0.0028638397,"teacher_disagreement_score":0.0069109653,"about_ca_system_score_codex":0.00035011384,"about_ca_system_score_gemma":0.0009844472,"threshold_uncertainty_score":0.02311951},"labels":[],"label_agreement":null},{"id":"W2478941700","doi":"10.1075/cilt.309.18isl","title":"Semantic similarity of short texts","year":2009,"lang":"en","type":"article","venue":"Amsterdam studies in the theory and history of linguistic science. Series 4, Current issues in linguistic theory","topic":"Topic Modeling","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Semantic similarity; Natural language processing; Similarity (geometry); Artificial intelligence; Word (group theory); Sentence; Variety (cybernetics); Focus (optics); Representation (politics); Matching (statistics); Longest common subsequence problem; Information retrieval; String metric; Pattern matching; String searching algorithm; Mathematics; Algorithm","score_opus":0.04907162359263313,"score_gpt":0.34327090327501936,"score_spread":0.29419927968238624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2478941700","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17840515,0.008344221,0.7790045,0.00089166604,0.0010222418,0.0007867875,0.0055169766,0.0015109108,0.024517527],"genre_scores_gemma":[0.62614524,0.0029299373,0.35386255,0.00028046494,0.0009608389,0.0011105172,0.008528199,0.00037611584,0.0058061252],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949091,0.0012581715,0.0007212418,0.0009807007,0.0020004727,0.00013038826],"domain_scores_gemma":[0.9892308,0.005699278,0.0014402304,0.0010532474,0.0022859103,0.0002906069],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022253077,0.0005744817,0.0010060746,0.013879479,0.0010010455,0.0028320786,0.0010416285,0.0010697315,0.0044196625],"category_scores_gemma":[0.026171364,0.00026205636,0.00084366265,0.008815708,0.0012745805,0.0063314834,0.0019703195,0.0008019678,0.0013417732],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010963042,0.0002685451,0.019925674,0.002661839,0.00065262715,0.0009727247,0.0055441177,0.015242181,0.03252784,0.19786257,0.011533982,0.7117116],"study_design_scores_gemma":[0.00014125148,0.0008547334,0.04867603,0.00092424225,0.00059787603,0.0034557052,0.006150487,0.19310455,0.027555503,0.5649153,0.15327436,0.00034994003],"about_ca_topic_score_codex":0.00085039635,"about_ca_topic_score_gemma":0.00080311915,"teacher_disagreement_score":0.013879479,"about_ca_system_score_codex":0.0009226873,"about_ca_system_score_gemma":0.00089087297,"threshold_uncertainty_score":0.01478523},"labels":[],"label_agreement":null},{"id":"W2481031298","doi":"10.4018/978-1-60566-274-9.ch011","title":"Analyzing the Text of Clinical Literature for Question Answering","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Ontario Institute for Cancer Research","funders":"","keywords":"Question answering; Computer science; Class (philosophy); Task (project management); Natural language processing; Information retrieval; Focus (optics); Artificial intelligence","score_opus":0.03597692082225227,"score_gpt":0.3213564867693303,"score_spread":0.28537956594707803,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2481031298","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022367159,0.06544061,0.77086055,0.011620878,0.0023449122,0.0018511963,0.017492238,0.008620462,0.09940187],"genre_scores_gemma":[0.08845947,0.024736749,0.8121421,0.0019511288,0.0018194205,0.0017279467,0.037068386,0.0018675675,0.030227209],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99795866,0.00087529246,0.00020149128,0.00037271337,0.00051595736,0.00007584954],"domain_scores_gemma":[0.9896376,0.00866247,0.00032252003,0.0005097957,0.00069494225,0.00017252359],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002859986,0.0017367797,0.0011608236,0.014127308,0.0013275272,0.005733108,0.0016863155,0.0015302664,0.021866193],"category_scores_gemma":[0.007981642,0.00060578604,0.0017922677,0.013430031,0.0012667128,0.00692399,0.001934194,0.0019165959,0.013568486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007767336,0.0001386649,0.0012291035,0.0038469255,0.000117507734,0.0005760758,0.004408741,0.0033616128,0.014788277,0.09796344,0.13627382,0.7372182],"study_design_scores_gemma":[0.000029599683,0.000078206576,0.0035706316,0.0020102214,0.00012194888,0.0013279297,0.002856207,0.03220954,0.005138288,0.12360007,0.82898086,0.00007655537],"about_ca_topic_score_codex":0.0013904937,"about_ca_topic_score_gemma":0.0019518278,"teacher_disagreement_score":0.021866193,"about_ca_system_score_codex":0.0020335258,"about_ca_system_score_gemma":0.0015602807,"threshold_uncertainty_score":0.07314974},"labels":[],"label_agreement":null},{"id":"W2481687795","doi":"10.4018/978-1-60566-908-3.ch004","title":"Concept-Based Mining Model","year":2010,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Sentence; Natural language processing; Artificial intelligence; Phrase; Term (time); Categorization; Representation (politics); Semantics (computer science); Graph; Information retrieval; Theoretical computer science","score_opus":0.029946248302209835,"score_gpt":0.24922540745103447,"score_spread":0.21927915914882462,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2481687795","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008560112,0.0014302523,0.9513323,0.0021501624,0.00019348886,0.0010643882,0.0028152156,0.00095006323,0.03150407],"genre_scores_gemma":[0.15281042,0.0023619062,0.817491,0.00082284055,0.00025810295,0.0018602527,0.005164322,0.00013996314,0.019091194],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976827,0.0005390936,0.00019304553,0.00062517176,0.0008541411,0.000105788495],"domain_scores_gemma":[0.99709964,0.001570904,0.00019705381,0.00026245505,0.00079189206,0.00007807373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024058565,0.0011672188,0.000985417,0.0034611267,0.000838336,0.0036134373,0.004425572,0.0013445773,0.012631536],"category_scores_gemma":[0.0070471726,0.0004330274,0.001902132,0.004330192,0.0009539842,0.005824092,0.0014179925,0.0015181943,0.005563284],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015527481,0.00032100704,0.0028465048,0.000914467,0.0002505558,0.0006319156,0.0008620142,0.049531948,0.0018626202,0.61866826,0.02792986,0.29602563],"study_design_scores_gemma":[0.00006987081,0.00008589045,0.0008166985,0.00018583646,0.00012210963,0.0009845882,0.00027582518,0.43868303,0.0014516563,0.4638289,0.093439296,0.00005630159],"about_ca_topic_score_codex":0.003936108,"about_ca_topic_score_gemma":0.002888036,"teacher_disagreement_score":0.012631536,"about_ca_system_score_codex":0.0016260403,"about_ca_system_score_gemma":0.0025001895,"threshold_uncertainty_score":0.042256713},"labels":[],"label_agreement":null},{"id":"W2483791419","doi":"10.1109/uic-atc-scalcom-cbdcom-iop.2015.163","title":"A Cloud Based Framework for Identification of Influential Health Experts from Twitter","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"","keywords":"Cloud computing; Computer science; Scalability; Identification (biology); Social media; Data science; Metric (unit); Health care; Hyperlink; World Wide Web; Big data; Data mining; Web page; Database; Engineering","score_opus":0.08587095000544487,"score_gpt":0.3423231894636827,"score_spread":0.2564522394582378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2483791419","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0437441,0.0007790757,0.9470734,0.00083499326,0.00010919193,0.0004841109,0.0011458712,0.0016592725,0.0041700564],"genre_scores_gemma":[0.5743822,0.00044106357,0.41869974,0.00018174406,0.00021147662,0.00031632045,0.0019712625,0.000085597785,0.0037105503],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988574,0.00028987057,0.000082775325,0.00026879788,0.0003171608,0.00018402236],"domain_scores_gemma":[0.9980572,0.00078154745,0.0002235999,0.00020886531,0.00052247936,0.00020637292],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015498113,0.00079983013,0.0010484114,0.002877965,0.0017151975,0.001957492,0.0015525188,0.0010583124,0.0020066795],"category_scores_gemma":[0.0039637797,0.00034541273,0.0010609471,0.0026393575,0.00042330995,0.002049009,0.0014645417,0.0006770742,0.0010324214],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015335393,0.0011235583,0.053729497,0.0005578014,0.00044875796,0.0019151444,0.0021255866,0.30622226,0.031309694,0.09159677,0.032203123,0.47723418],"study_design_scores_gemma":[0.000013251835,0.000018448143,0.0009479253,0.000007570843,0.000019936833,0.00007121416,0.00015516143,0.99090284,0.001124006,0.0045461096,0.0021809526,0.000012528171],"about_ca_topic_score_codex":0.030890761,"about_ca_topic_score_gemma":0.04472511,"teacher_disagreement_score":0.030890761,"about_ca_system_score_codex":0.0014320797,"about_ca_system_score_gemma":0.0025543058,"threshold_uncertainty_score":0.06142193},"labels":[],"label_agreement":null},{"id":"W2484517919","doi":"10.4018/978-1-60960-741-8.ch013","title":"HiDEx","year":2012,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Representation (politics); Computer science; Hyperspace; Set (abstract data type); Basis (linear algebra); A priori and a posteriori; Semantic space; Space (punctuation); Natural language processing; Artificial intelligence; Matrix (chemical analysis); Mathematics; Programming language; Epistemology","score_opus":0.02968519643606811,"score_gpt":0.2441832130512419,"score_spread":0.21449801661517381,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2484517919","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010383951,0.019445218,0.42771828,0.0049366066,0.0011399791,0.0001234383,0.0049834307,0.0068443217,0.52442473],"genre_scores_gemma":[0.17302881,0.028586011,0.20478491,0.0021576895,0.0009615207,0.00046326494,0.008589243,0.003143848,0.5782847],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998839,0.000027485317,0.0000058709284,0.0000341386,0.000041385196,0.000007250037],"domain_scores_gemma":[0.99981564,0.000102403894,0.000010690379,0.000040850744,0.000016626573,0.000013816235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00020800899,0.000689882,0.00035846705,0.0005821549,0.00036508744,0.0019487304,0.00076551404,0.00066900614,0.046809837],"category_scores_gemma":[0.0007546142,0.0002698144,0.0004486159,0.00093482045,0.0008361712,0.0034887781,0.00114209,0.0012580635,0.015278738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037721464,0.0000133933745,0.00032658144,0.0002894698,0.000020658857,0.000075202464,0.00036712005,0.004014374,0.0013576141,0.7363549,0.090086296,0.16705666],"study_design_scores_gemma":[0.000008263409,0.000017785243,0.00037440623,0.00009092919,0.000008863423,0.00022239114,0.000077224926,0.010042633,0.0007624477,0.21871468,0.76966584,0.000014421402],"about_ca_topic_score_codex":0.0013309334,"about_ca_topic_score_gemma":0.0016282178,"teacher_disagreement_score":0.046809837,"about_ca_system_score_codex":0.0006183936,"about_ca_system_score_gemma":0.00043693435,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2489491723","doi":"10.1007/978-3-319-41754-7_8","title":"Evaluating Multiple Summaries Without Human Models: A First Experiment with a Trivergent Model","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Task (project management); Set (abstract data type); Natural language processing; Artificial intelligence; Information retrieval; Source model; Programming language; Theoretical computer science; Systems engineering","score_opus":0.08502014523120585,"score_gpt":0.3093159256192198,"score_spread":0.22429578038801393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2489491723","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8976561,0.0060527264,0.06965659,0.0019983782,0.0014915986,0.0018621206,0.0055316747,0.009315999,0.0064348285],"genre_scores_gemma":[0.8481977,0.001144008,0.12868771,0.00090065825,0.00040784793,0.00079750764,0.012740892,0.0010240367,0.00609952],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98913974,0.006232912,0.0010894787,0.0023438805,0.0009921888,0.0002017983],"domain_scores_gemma":[0.88383657,0.09829517,0.0017864103,0.0094875265,0.003725401,0.0028689615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014208421,0.002550972,0.0028896444,0.0014524009,0.0016163775,0.0041875597,0.0039252895,0.005463657,0.007610294],"category_scores_gemma":[0.07922464,0.0009881986,0.0018187924,0.0015423673,0.0008139679,0.0074687265,0.0020774286,0.0039988155,0.0030514507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.049374208,0.01931606,0.024618564,0.008785616,0.005582627,0.0022573937,0.008336811,0.15210402,0.055690672,0.004821812,0.07579068,0.5933216],"study_design_scores_gemma":[0.0073294514,0.02537986,0.01239416,0.0004235327,0.0027763348,0.0013840293,0.0021667334,0.89274484,0.021471435,0.008260973,0.025141243,0.0005274968],"about_ca_topic_score_codex":0.007970866,"about_ca_topic_score_gemma":0.008521074,"teacher_disagreement_score":0.014208421,"about_ca_system_score_codex":0.0012272059,"about_ca_system_score_gemma":0.001746225,"threshold_uncertainty_score":0.075142205},"labels":[],"label_agreement":null},{"id":"W2490635148","doi":"","title":"Automatic Question Generation from Sentences","year":2010,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Sentence; Natural language processing; Question answering; Artificial intelligence; Parsing; Task (project management); Set (abstract data type); Subject (documents); Natural language; Linguistics; World Wide Web; Programming language","score_opus":0.022603715257053743,"score_gpt":0.2544485896510814,"score_spread":0.23184487439402768,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2490635148","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05348274,0.0017894247,0.86407375,0.0016281497,0.00071793504,0.0019090374,0.02540129,0.039840955,0.011156672],"genre_scores_gemma":[0.20037189,0.0007891133,0.7026949,0.0005155614,0.0003445905,0.0012598239,0.085454985,0.0020790633,0.0064900415],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967834,0.0013938855,0.00026654472,0.0007760598,0.0005990631,0.00018099743],"domain_scores_gemma":[0.99178773,0.004721898,0.0003840897,0.0008449495,0.0020763,0.00018509885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024356954,0.0015878191,0.0010585291,0.0028260455,0.00080594193,0.0015437528,0.0013122689,0.0013603807,0.011114103],"category_scores_gemma":[0.010952447,0.00060471456,0.0016264437,0.0018149277,0.0004624938,0.0020767692,0.0019174726,0.0012433289,0.0065030064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081428804,0.00038080016,0.0052379137,0.002043641,0.00020019815,0.0013232472,0.0019329521,0.01273937,0.09101057,0.030220546,0.15556036,0.6985361],"study_design_scores_gemma":[0.00042595656,0.0004910408,0.008888502,0.00041020053,0.00033797798,0.0021655946,0.0016481151,0.5130871,0.11954331,0.10982609,0.24296032,0.0002158377],"about_ca_topic_score_codex":0.0016075907,"about_ca_topic_score_gemma":0.0016537595,"teacher_disagreement_score":0.011114103,"about_ca_system_score_codex":0.0008482355,"about_ca_system_score_gemma":0.0015604456,"threshold_uncertainty_score":0.037180364},"labels":[],"label_agreement":null},{"id":"W2509318466","doi":"10.18653/v1/w16-0302","title":"Towards Early Dementia Detection: Fusing Linguistic and Non-Linguistic Clinical Data","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute on Aging; National Institute of Biomedical Imaging and Bioengineering; Canadian Institutes of Health Research; National Institutes of Health; Genentech; IXICO; Servier; Eisai; Northern California Institute for Research and Education; University of California, San Diego; Pfizer; Biogen; BioClinica; Eli Lilly and Company; U.S. Department of Defense; Meso Scale Diagnostics; Synarc; University of Southern California; Medpace; Novartis Pharmaceuticals Corporation; Bristol-Myers Squibb; F. Hoffmann-La Roche; Alzheimer's Drug Discovery Foundation; Foundation for the National Institutes of Health","keywords":"Computer science; Dementia; Linguistics; Natural language processing; Linguistic analysis; Artificial intelligence; Deep linguistic processing; Medicine; Disease; Philosophy","score_opus":0.09277553604152745,"score_gpt":0.34425139838408464,"score_spread":0.2514758623425572,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2509318466","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3997503,0.0111306105,0.55928284,0.006779139,0.0005192118,0.0008242477,0.008396269,0.004833509,0.008483923],"genre_scores_gemma":[0.78148174,0.0022889522,0.203852,0.0005332655,0.00028226347,0.00030015904,0.009448946,0.0001574952,0.0016551713],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953442,0.0029203957,0.00040350185,0.00072191824,0.00039235962,0.000217523],"domain_scores_gemma":[0.9733376,0.021753743,0.0014193107,0.0013053425,0.0017556308,0.00042835035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00973621,0.0016953424,0.0012700967,0.008709177,0.0006568729,0.0037565841,0.001057559,0.0013504584,0.0014979359],"category_scores_gemma":[0.03624082,0.00051374733,0.001532445,0.004495447,0.00054194673,0.0051910346,0.002830821,0.0017875179,0.0012761573],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016786451,0.0010858165,0.13234103,0.0016509746,0.0009226751,0.0007406278,0.0035316416,0.020424306,0.014648958,0.003116893,0.00817852,0.81167984],"study_design_scores_gemma":[0.00015504986,0.00092244527,0.11683186,0.001268254,0.0017742657,0.0012397857,0.007007772,0.76239336,0.014800793,0.07036334,0.022854866,0.0003882412],"about_ca_topic_score_codex":0.0067027872,"about_ca_topic_score_gemma":0.008479552,"teacher_disagreement_score":0.00973621,"about_ca_system_score_codex":0.00069947384,"about_ca_system_score_gemma":0.0016271848,"threshold_uncertainty_score":0.051490605},"labels":[],"label_agreement":null},{"id":"W2510135622","doi":"10.48550/arxiv.1608.07738","title":"Testing APSyn against Vector Cosine on Similarity Estimation","year":2016,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited; University of Oxford; University of Wisconsin-Madison","keywords":"Cosine similarity; Similarity (geometry); Computer science; Discrete cosine transform; Estimation; Artificial intelligence; Pattern recognition (psychology); Algorithm; Engineering","score_opus":0.11816236527354619,"score_gpt":0.20107105186456534,"score_spread":0.08290868659101915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2510135622","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7846924,0.012706967,0.16563612,0.002057095,0.0020525742,0.0007347712,0.0065502683,0.005645688,0.019924069],"genre_scores_gemma":[0.92507696,0.00073369354,0.06308024,0.00037591267,0.0003892992,0.0003130554,0.008040664,0.00035879042,0.001631456],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9599506,0.023594847,0.0028726126,0.0055389595,0.0070869206,0.00095591805],"domain_scores_gemma":[0.89335036,0.08412988,0.0032651364,0.011673861,0.005799671,0.0017811194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031130135,0.002516489,0.0024269742,0.005712111,0.0014159476,0.0034359645,0.0033785533,0.00360436,0.004499046],"category_scores_gemma":[0.12988667,0.00048150416,0.0014352655,0.005397126,0.0023538677,0.0105427215,0.0060979696,0.0024698453,0.0027907926],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009833773,0.0022158604,0.1425074,0.002805374,0.0039624604,0.00044814125,0.0008289262,0.13872543,0.00650275,0.02127917,0.036662623,0.6342281],"study_design_scores_gemma":[0.0005303957,0.004126378,0.024097554,0.00022584703,0.0004381413,0.0009513169,0.0012480365,0.9313679,0.007305636,0.022222742,0.007339782,0.00014618498],"about_ca_topic_score_codex":0.0039878553,"about_ca_topic_score_gemma":0.0029577538,"teacher_disagreement_score":0.031130135,"about_ca_system_score_codex":0.0011589415,"about_ca_system_score_gemma":0.0017099238,"threshold_uncertainty_score":0.16463387},"labels":[],"label_agreement":null},{"id":"W2510461200","doi":"10.18653/v1/w16-0529","title":"Combining Off-the-shelf Grammar and Spelling Tools for the Automatic Evaluation of Scientific Writing (AESW) Shared Task 2016","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Stroke Network","funders":"","keywords":"Spelling; Computer science; Task (project management); Grammar; Natural language processing; Sentence; Artificial intelligence; Recall; Binary number; Binary classification; Linguistics; Arithmetic","score_opus":0.08630704066969959,"score_gpt":0.29780336166516663,"score_spread":0.21149632099546706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2510461200","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46687847,0.0021728475,0.30375805,0.0013389416,0.0020530557,0.003752708,0.03702114,0.1689503,0.014074501],"genre_scores_gemma":[0.6147305,0.00036475842,0.2968578,0.0005271521,0.00040727895,0.0033391824,0.06634565,0.010503992,0.006923738],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97743326,0.009461634,0.0038377603,0.0043144263,0.0040702596,0.0008826532],"domain_scores_gemma":[0.9438443,0.026864212,0.0031225614,0.0085430825,0.015416925,0.0022089768],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016530262,0.003488813,0.0020621477,0.0073075253,0.001122915,0.0030399878,0.0024751225,0.0028149972,0.005092372],"category_scores_gemma":[0.06146045,0.0010062763,0.0018557814,0.0021129397,0.0007798352,0.0037949025,0.0068290825,0.0024426917,0.008282926],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014896555,0.0015341011,0.020099575,0.0023690574,0.00067033025,0.0010251016,0.0028179463,0.011347447,0.084666945,0.0011877101,0.09406548,0.7787267],"study_design_scores_gemma":[0.0019662862,0.0033124995,0.1082899,0.000798786,0.0006680533,0.0030835134,0.0032282355,0.49267316,0.25247058,0.015881455,0.11657233,0.0010553261],"about_ca_topic_score_codex":0.004435483,"about_ca_topic_score_gemma":0.005710825,"teacher_disagreement_score":0.016530262,"about_ca_system_score_codex":0.0012318785,"about_ca_system_score_gemma":0.0032816571,"threshold_uncertainty_score":0.08742142},"labels":[],"label_agreement":null},{"id":"W2510939506","doi":"10.18653/v1/w16-0417","title":"Semi-supervised and unsupervised categorization of posts in Web discussion forums using part-of-speech information and minimal features","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Categorization; Artificial intelligence; Cluster analysis; Hidden Markov model; Identification (biology); Probabilistic logic; Topic model; Natural language processing; Machine learning; Unsupervised learning; Information retrieval","score_opus":0.0154227995785618,"score_gpt":0.23465858357616196,"score_spread":0.21923578399760016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2510939506","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27135932,0.00087717036,0.7160844,0.00026062122,0.00014388857,0.0007233865,0.0017624293,0.004151946,0.0046368376],"genre_scores_gemma":[0.7596779,0.0002754341,0.2278784,0.00009550573,0.00019260056,0.0005858872,0.006375972,0.00024457907,0.004673709],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99632335,0.0016028641,0.0002357803,0.0011071841,0.0005383478,0.0001924469],"domain_scores_gemma":[0.9913788,0.0047484445,0.0011050376,0.0010059545,0.0014718471,0.00029000785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035521379,0.0012038822,0.0010484179,0.0048511797,0.000838189,0.0011857976,0.0015367599,0.0010507966,0.0010897757],"category_scores_gemma":[0.008339812,0.00039953328,0.001292624,0.0017061513,0.00078857504,0.0023550976,0.00097217987,0.0009437289,0.0019793005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011734454,0.0016200929,0.040420678,0.0011968297,0.0005475759,0.00024575577,0.00228656,0.03879073,0.046347525,0.004516859,0.012320581,0.85053337],"study_design_scores_gemma":[0.000068915586,0.0003781526,0.028718589,0.00010376302,0.00014256114,0.00036257633,0.0008876164,0.93034774,0.023949565,0.009867778,0.005045996,0.00012664907],"about_ca_topic_score_codex":0.0019593644,"about_ca_topic_score_gemma":0.0052197766,"teacher_disagreement_score":0.0048511797,"about_ca_system_score_codex":0.00065616146,"about_ca_system_score_gemma":0.0013932235,"threshold_uncertainty_score":0.018785715},"labels":[],"label_agreement":null},{"id":"W2511646563","doi":"10.18653/v1/p16-1108","title":"Leveraging Inflection Tables for Stemming and Lemmatization.","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Lemmatisation; Inflection; Computer science; Artificial intelligence; Discriminative model; Natural language processing; String (physics); Task (project management); Exploit; Mathematics","score_opus":0.03247729577101645,"score_gpt":0.24366172841010153,"score_spread":0.2111844326390851,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2511646563","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012588036,0.00060190965,0.95636606,0.00026118386,0.00015770816,0.00016444564,0.0029184942,0.0235965,0.0033455964],"genre_scores_gemma":[0.15502994,0.0005613522,0.81517494,0.00026360629,0.00015705533,0.00019727767,0.021983838,0.0024888257,0.004143229],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99812084,0.00059440034,0.00021222669,0.0006302625,0.0003485003,0.00009365698],"domain_scores_gemma":[0.9941229,0.0028436591,0.00046388866,0.0015998139,0.0008405011,0.00012935606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023733466,0.001609302,0.0006237027,0.0043384214,0.00082086795,0.0024903624,0.0015041814,0.0013213948,0.0055802325],"category_scores_gemma":[0.009915613,0.0006628863,0.0014463608,0.0035479954,0.000696232,0.00479896,0.0021627613,0.0020336602,0.013427901],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002091348,0.00019321288,0.006470235,0.00077434693,0.0002458485,0.0003567175,0.0008902212,0.012317234,0.047843333,0.019290255,0.037314806,0.8740948],"study_design_scores_gemma":[0.000091869944,0.00022416847,0.00839527,0.00021871015,0.00020902015,0.001920513,0.00083519914,0.6424061,0.12201251,0.096514806,0.12698139,0.00019048023],"about_ca_topic_score_codex":0.0014357925,"about_ca_topic_score_gemma":0.004129004,"teacher_disagreement_score":0.0055802325,"about_ca_system_score_codex":0.0005106323,"about_ca_system_score_gemma":0.0014384901,"threshold_uncertainty_score":0.018667758},"labels":[],"label_agreement":null},{"id":"W2512683849","doi":"10.18653/v1/w16-0414","title":"Classification of comment helpfulness to improve knowledge sharing among medical practitioners.","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Computer Research Institute of Montréal","funders":"","keywords":"Helpfulness; Computer science; Information retrieval; Relevance (law); Binary classification; Baseline (sea); Natural language processing; Artificial intelligence; Support vector machine; Psychology","score_opus":0.03852776795705148,"score_gpt":0.2995865932101967,"score_spread":0.2610588252531452,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2512683849","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9246606,0.0038972711,0.05105405,0.0012900649,0.00034753905,0.0011864109,0.008260002,0.0030168858,0.0062870462],"genre_scores_gemma":[0.93091476,0.00035256837,0.056045286,0.00012280387,0.000317926,0.00046226865,0.009923707,0.000096740274,0.0017639098],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913783,0.0042141587,0.0008400532,0.0012600226,0.001953674,0.00035380936],"domain_scores_gemma":[0.92070395,0.05435428,0.007025566,0.002728797,0.013326423,0.0018610925],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.010231441,0.0010223376,0.00079954934,0.009490643,0.0008390449,0.0015184105,0.0009873309,0.0014784378,0.0013195417],"category_scores_gemma":[0.05890797,0.00021112569,0.00077748566,0.0032701634,0.0003858145,0.0029456178,0.0015873332,0.0010076714,0.0010316048],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003229779,0.0019455808,0.31288862,0.003143917,0.0006636047,0.0005816577,0.004907016,0.005680571,0.031618435,0.0011315328,0.030861795,0.60334754],"study_design_scores_gemma":[0.0003859879,0.0031157942,0.51138294,0.00066427304,0.0010063121,0.0013334972,0.0073206536,0.4096523,0.03200855,0.0061431816,0.02670773,0.000278885],"about_ca_topic_score_codex":0.0028225484,"about_ca_topic_score_gemma":0.0054070256,"teacher_disagreement_score":0.9984816,"about_ca_system_score_codex":0.00092736026,"about_ca_system_score_gemma":0.0012263515,"threshold_uncertainty_score":0.054109633},"labels":[],"label_agreement":null},{"id":"W2514692010","doi":"10.18653/v1/p16-1221","title":"Vector-space topic models for detecting Alzheimer's disease","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Toronto Rehabilitation Institute","funders":"Toronto Rehabilitation Institute","keywords":"Computer science; Space (punctuation); Vector (molecular biology); Disease; Artificial intelligence; Medicine; Biology","score_opus":0.05309899302884141,"score_gpt":0.2678546956129835,"score_spread":0.21475570258414206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2514692010","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.086978,0.0034810258,0.9036818,0.00040513472,0.00014317948,0.00023065186,0.001147063,0.0028191428,0.0011140419],"genre_scores_gemma":[0.71588904,0.001851348,0.270934,0.00016348742,0.00047113234,0.0005614874,0.005937892,0.00032088533,0.0038706677],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987985,0.00060000585,0.00008854721,0.00024865812,0.00017654712,0.00008771857],"domain_scores_gemma":[0.99574983,0.0034648383,0.00020995195,0.00017901405,0.00033557438,0.0000606932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033612205,0.0013145631,0.0008287773,0.00309371,0.0004971696,0.0011566966,0.0010967609,0.001102535,0.0012054592],"category_scores_gemma":[0.0065006837,0.00044701897,0.0013104571,0.0018063079,0.00041230593,0.0015226593,0.0007803429,0.0013923236,0.0008390741],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074744096,0.0003547476,0.013038506,0.0005356447,0.0005878176,0.00024373052,0.0007957424,0.40046772,0.012845305,0.017281849,0.011404398,0.5416971],"study_design_scores_gemma":[0.00002730619,0.00006521238,0.0015276625,0.000019337522,0.000052836484,0.000066640656,0.000043640448,0.9863347,0.001548717,0.008617994,0.001678278,0.000017663911],"about_ca_topic_score_codex":0.0063426653,"about_ca_topic_score_gemma":0.005611058,"teacher_disagreement_score":0.0063426653,"about_ca_system_score_codex":0.0007460745,"about_ca_system_score_gemma":0.0006347086,"threshold_uncertainty_score":0.017776072},"labels":[],"label_agreement":null},{"id":"W2516087440","doi":"10.18653/v1/p16-1063","title":"Generative Topic Embedding: a Continuous Representation of Documents","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":108,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Research (Canada)","funders":"National Research Foundation","keywords":"Embedding; Computer science; Representation (politics); Generative grammar; Natural language processing; Artificial intelligence; Political science","score_opus":0.028558937665621217,"score_gpt":0.3116345904051909,"score_spread":0.2830756527395697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2516087440","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006064392,0.0006423066,0.99132794,0.00022047348,0.000052697844,0.00004046978,0.0004430599,0.00052170816,0.0006869278],"genre_scores_gemma":[0.41804594,0.0025487307,0.56686884,0.00027297367,0.00045452832,0.00054648257,0.0043824906,0.0005166257,0.0063633192],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99884194,0.0004441382,0.00006633357,0.00037030815,0.0001981677,0.00007912559],"domain_scores_gemma":[0.9976763,0.0014327377,0.0002062025,0.00035824915,0.00025427836,0.000072167175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015576242,0.0009410367,0.0007485375,0.0025023536,0.00041656895,0.0019424196,0.001490668,0.0011439715,0.0024783965],"category_scores_gemma":[0.006944601,0.00053082185,0.0012290581,0.0030857683,0.00076814904,0.0032701169,0.0013274786,0.0018230425,0.0010305798],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034847495,0.00019092111,0.0047181337,0.000634036,0.00034322403,0.00035674713,0.0014327717,0.24091662,0.012120478,0.18964,0.019195072,0.53010345],"study_design_scores_gemma":[0.0000240801,0.000038965383,0.0008792675,0.000043378546,0.000039189887,0.00014404573,0.00007621439,0.91526824,0.0015769857,0.0744343,0.0074419244,0.000033393226],"about_ca_topic_score_codex":0.0031910555,"about_ca_topic_score_gemma":0.0033414091,"teacher_disagreement_score":0.0031910555,"about_ca_system_score_codex":0.00078745204,"about_ca_system_score_gemma":0.0008999947,"threshold_uncertainty_score":0.008291066},"labels":[],"label_agreement":null},{"id":"W2517784737","doi":"10.18653/v1/k16-1024","title":"Event Linking with Sentential Features from Convolutional Neural Networks","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Coreference; Computer science; Pairwise comparison; Artificial intelligence; Event (particle physics); Convolutional neural network; Context (archaeology); Natural language processing; Similarity (geometry); Feature (linguistics); Task (project management); Process (computing); Machine learning; Resolution (logic); Image (mathematics)","score_opus":0.011228276694675243,"score_gpt":0.2121893518677359,"score_spread":0.20096107517306067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2517784737","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11123043,0.0014408055,0.87123376,0.00089776254,0.00015748419,0.000147884,0.0010420645,0.008412102,0.0054377844],"genre_scores_gemma":[0.8125041,0.00071712077,0.17287397,0.00029185676,0.00016773629,0.00023430673,0.003806746,0.00033516876,0.009068928],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995466,0.00009484929,0.000029317218,0.00018390609,0.0000718275,0.000073416835],"domain_scores_gemma":[0.9989405,0.0005755951,0.000116418436,0.00018991661,0.00014212314,0.00003539908],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013370852,0.001401413,0.0006759301,0.0017401553,0.00056896976,0.0013195951,0.0018897131,0.0014708177,0.0021891887],"category_scores_gemma":[0.0038409084,0.0006688958,0.0010802874,0.0018919171,0.0005389081,0.003417646,0.0016743763,0.0022665467,0.0011175398],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043244858,0.00028291138,0.0036317762,0.00019499849,0.00024034845,0.00025750237,0.0002870089,0.3896172,0.013990402,0.014061765,0.01025077,0.56675285],"study_design_scores_gemma":[0.0000103716375,0.00001795156,0.00045230976,0.000012472735,0.00002918662,0.00002886648,0.0000142172885,0.9829401,0.0034868023,0.011832216,0.0011662124,0.000009308262],"about_ca_topic_score_codex":0.008591551,"about_ca_topic_score_gemma":0.016495747,"teacher_disagreement_score":0.008591551,"about_ca_system_score_codex":0.0015476629,"about_ca_system_score_gemma":0.00087150524,"threshold_uncertainty_score":0.017083108},"labels":[],"label_agreement":null},{"id":"W2524313288","doi":"10.21700/ijcis.2016.109","title":"Textual Entailment for Arabic Language based on Lexical and Semantic Matching","year":2016,"lang":"en","type":"article","venue":"International Journal of Computing and Information Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Natural language processing; Arabic; Computer science; Artificial intelligence; Textual entailment; Logical consequence; Matching (statistics); Linguistics; Mathematics; Philosophy","score_opus":0.016973547757490205,"score_gpt":0.30743686356854893,"score_spread":0.2904633158110587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2524313288","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14360268,0.0016335952,0.8312115,0.0014347595,0.0003571436,0.00063112413,0.0052514216,0.004208279,0.011669564],"genre_scores_gemma":[0.61569583,0.00078403,0.3646902,0.00020889241,0.00032575862,0.00034200554,0.01127713,0.00043329914,0.006242879],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99753964,0.0006926799,0.00034369732,0.0006054998,0.0006274873,0.00019103877],"domain_scores_gemma":[0.99691004,0.0014499995,0.00023421542,0.00036166285,0.0009512966,0.00009279982],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013446657,0.0006752441,0.0008842128,0.0033266465,0.0014518843,0.0021672258,0.0011373931,0.0010208169,0.010481028],"category_scores_gemma":[0.007557885,0.00035402857,0.0016702175,0.0021684852,0.00059025397,0.0044546006,0.0016140676,0.0010489863,0.0030696313],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022688555,0.0007609849,0.0061385827,0.0019721503,0.00036747256,0.0019088404,0.0016459689,0.018243818,0.062964804,0.10005051,0.03145485,0.7722232],"study_design_scores_gemma":[0.00022813167,0.00036221245,0.0064122397,0.00022753589,0.0006805384,0.0016982711,0.0017287342,0.7562686,0.058326513,0.14593273,0.027956938,0.00017755361],"about_ca_topic_score_codex":0.0037371577,"about_ca_topic_score_gemma":0.0035676397,"teacher_disagreement_score":0.010481028,"about_ca_system_score_codex":0.00073177397,"about_ca_system_score_gemma":0.0018112201,"threshold_uncertainty_score":0.03506255},"labels":[],"label_agreement":null},{"id":"W2524791359","doi":"10.1007/s11192-016-2134-8","title":"The linguistic patterns and rhetorical structure of citation context: an approach using n-grams","year":2016,"lang":"en","type":"article","venue":"Scientometrics","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Rhetorical question; Citation; Linguistics; Context (archaeology); Computer science; Linguistic context; Function (biology); Natural (archaeology); Natural language processing; Information retrieval; Linguistic analysis; History; Library science; Philosophy; Biology","score_opus":0.08286568018633636,"score_gpt":0.31700194935228493,"score_spread":0.23413626916594857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2524791359","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3768719,0.006962447,0.5785939,0.0026854863,0.0003745612,0.00074577396,0.007578831,0.0024654223,0.023721654],"genre_scores_gemma":[0.83181566,0.0016836083,0.15940082,0.00014144422,0.00040639684,0.00062012934,0.0029229422,0.00031841558,0.0026906165],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99446565,0.0029269957,0.0005187229,0.0007337331,0.0011923534,0.0001624626],"domain_scores_gemma":[0.958456,0.03586601,0.0023588357,0.001162005,0.0018167933,0.00034040652],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.006477996,0.0009318867,0.0012683982,0.02799213,0.0019810512,0.005871946,0.0012719424,0.0017523484,0.0023252412],"category_scores_gemma":[0.039522123,0.0006491996,0.0012569443,0.02988831,0.0013966445,0.007938967,0.0019417485,0.0015722293,0.0009922113],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009291704,0.00088047463,0.11300619,0.0026026124,0.0012989731,0.00082550436,0.0123678725,0.023675742,0.019417124,0.16183953,0.010990588,0.6521662],"study_design_scores_gemma":[0.00011022818,0.0002922655,0.057552304,0.0007052216,0.0011027681,0.0007625041,0.0051315357,0.5478707,0.008416861,0.341279,0.036535654,0.00024100214],"about_ca_topic_score_codex":0.0033211506,"about_ca_topic_score_gemma":0.0064835465,"teacher_disagreement_score":0.993522,"about_ca_system_score_codex":0.0015032729,"about_ca_system_score_gemma":0.0021792578,"threshold_uncertainty_score":0.03425932},"labels":[],"label_agreement":null},{"id":"W2530085701","doi":"10.1007/978-3-319-48051-0_12","title":"Constraining Word Embeddings by Prior Knowledge – Application to Medical Information Retrieval","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Leverage (statistics); Word (group theory); Natural language processing; Word embedding; Artificial intelligence; Embedding; Domain (mathematical analysis); Information retrieval; Linguistics; Mathematics","score_opus":0.011942982021409365,"score_gpt":0.26223426959109486,"score_spread":0.2502912875696855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2530085701","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035477713,0.002398948,0.9567129,0.0006443726,0.00017449514,0.00009204434,0.0005587158,0.0019852638,0.0019555262],"genre_scores_gemma":[0.4054471,0.002834352,0.57805765,0.00040397432,0.0006054967,0.00027365214,0.0044827727,0.00089951995,0.0069954507],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.998611,0.0005163926,0.00015267011,0.00040431763,0.00021921354,0.00009653618],"domain_scores_gemma":[0.9930173,0.0049763895,0.00033023008,0.0007460467,0.00078942423,0.00014051919],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027770041,0.0013182077,0.0015004049,0.0026280717,0.0007032456,0.0022009737,0.0016626139,0.0020306546,0.0033884996],"category_scores_gemma":[0.013717867,0.0010799057,0.0014075388,0.0037622508,0.0009166686,0.0054311296,0.0027512005,0.0025962132,0.002093414],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004416301,0.00034433274,0.0015973601,0.00042193662,0.00016505815,0.00016534253,0.00029286806,0.11889235,0.007075059,0.010970625,0.012128051,0.84750545],"study_design_scores_gemma":[0.00004540977,0.00007528749,0.000740643,0.00005499962,0.000059364676,0.0001329918,0.000086877706,0.95348114,0.0021099048,0.040301397,0.002878423,0.00003352697],"about_ca_topic_score_codex":0.005444762,"about_ca_topic_score_gemma":0.0062483554,"teacher_disagreement_score":0.005444762,"about_ca_system_score_codex":0.00061958,"about_ca_system_score_gemma":0.0010369434,"threshold_uncertainty_score":0.014686406},"labels":[],"label_agreement":null},{"id":"W2530152826","doi":"10.1145/2872518.2889397","title":"A Machine learning Filter for Relation Extraction","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Polytechnique Montréal","funders":"","keywords":"Computer science; Relation (database); Relationship extraction; Extraction (chemistry); Filter (signal processing); Artificial intelligence; Machine learning; Data mining; Computer vision; Chromatography","score_opus":0.03431274194028436,"score_gpt":0.2715054275684631,"score_spread":0.23719268562817872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2530152826","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04337211,0.0014527795,0.9216021,0.00067138276,0.0003261298,0.00048514898,0.0037223515,0.025923302,0.0024446163],"genre_scores_gemma":[0.16423565,0.0004260016,0.8164495,0.00048565524,0.00026997886,0.0005660616,0.0091518,0.0007988673,0.0076164473],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99582714,0.0006115429,0.00051111623,0.0014001788,0.0012991941,0.00035077703],"domain_scores_gemma":[0.98930454,0.0059859087,0.0006197849,0.0011797574,0.0027248897,0.00018521986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057292553,0.0018069167,0.0019613777,0.006858823,0.0022399055,0.0028577242,0.002093416,0.002848552,0.0043975585],"category_scores_gemma":[0.011774852,0.00068870944,0.0020230524,0.0042729103,0.0006932433,0.0034585637,0.0013940411,0.0021088968,0.0051780054],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000581982,0.0003873927,0.0091685355,0.00052316894,0.00026173805,0.00044120778,0.00047334656,0.005708155,0.047428526,0.004559679,0.02869168,0.9017745],"study_design_scores_gemma":[0.0001602479,0.00066467543,0.0177485,0.00023688683,0.00055539794,0.0018241284,0.00045725497,0.6528383,0.20727186,0.01436914,0.10367788,0.00019577367],"about_ca_topic_score_codex":0.011081599,"about_ca_topic_score_gemma":0.012305862,"teacher_disagreement_score":0.011081599,"about_ca_system_score_codex":0.0015677179,"about_ca_system_score_gemma":0.002830342,"threshold_uncertainty_score":0.030299604},"labels":[],"label_agreement":null},{"id":"W2534147738","doi":"10.1145/2983323.2983694","title":"Optimizing Nugget Annotations with Active Learning","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; University of Waterloo; Google","keywords":"Automatic summarization; Computer science; Annotation; Sentence; Process (computing); Sequence (biology); Precision and recall; Recall; Natural language processing; Artificial intelligence; Information retrieval; Programming language","score_opus":0.017462470091850278,"score_gpt":0.236603102277309,"score_spread":0.2191406321854587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2534147738","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04505194,0.0013354382,0.9361083,0.0005588782,0.00021218196,0.00033701386,0.0006176841,0.011494288,0.004284277],"genre_scores_gemma":[0.5086246,0.00046785176,0.47109994,0.0007542155,0.00042432436,0.0006737875,0.0038998562,0.0015473211,0.012508108],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99497813,0.002107012,0.0002987205,0.0013618165,0.0008960312,0.00035837674],"domain_scores_gemma":[0.9774914,0.016560823,0.00093075196,0.0022549785,0.0022543573,0.00050774607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007440864,0.002638204,0.002217495,0.0033992007,0.0016475894,0.0029403078,0.00471112,0.0039273296,0.0045160395],"category_scores_gemma":[0.027403882,0.0010026916,0.0013449808,0.0026727987,0.0012022434,0.007225139,0.0033120953,0.0040970272,0.0026967276],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012996778,0.0010145083,0.00480748,0.00040837907,0.00022011161,0.00019531613,0.00096979894,0.18338102,0.013951036,0.007847095,0.0222865,0.7636191],"study_design_scores_gemma":[0.00007486174,0.00014975155,0.00042795323,0.00003217568,0.00006530353,0.00006191583,0.00013572605,0.97972804,0.0062816925,0.009799549,0.003214474,0.00002854223],"about_ca_topic_score_codex":0.00803819,"about_ca_topic_score_gemma":0.014560236,"teacher_disagreement_score":0.00803819,"about_ca_system_score_codex":0.0016703209,"about_ca_system_score_gemma":0.0018819914,"threshold_uncertainty_score":0.039351523},"labels":[],"label_agreement":null},{"id":"W2535516716","doi":"10.1109/wcse.2013.17","title":"Evaluation of Stability and Similarity of Latent Dirichlet Allocation","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Pivotal (Canada)","funders":"","keywords":"Latent Dirichlet allocation; Divergence (linguistics); Computer science; Similarity (geometry); Stability (learning theory); Categorization; Artificial intelligence; Topic model; Matching (statistics); Set (abstract data type); Pattern recognition (psychology); Probabilistic latent semantic analysis; Dirichlet distribution; Kullback–Leibler divergence; Key (lock); Machine learning; Data mining; Mathematics; Image (mathematics); Statistics","score_opus":0.09114239174412617,"score_gpt":0.2913854160830547,"score_spread":0.20024302433892852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2535516716","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5191549,0.0013516976,0.47334844,0.00040397551,0.00013483662,0.00032894465,0.00063746516,0.0014090782,0.0032306425],"genre_scores_gemma":[0.88523954,0.00015770145,0.11222804,0.00005317891,0.000056069162,0.00019122995,0.0013755255,0.00020170785,0.00049699214],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98468405,0.0069242786,0.0016789963,0.0024986,0.0037762148,0.0004378437],"domain_scores_gemma":[0.9398124,0.038839705,0.003754491,0.0066571464,0.009942033,0.0009941857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020706441,0.0010679624,0.0015112875,0.005223861,0.0012468677,0.0027259912,0.0015046871,0.002045218,0.00071151304],"category_scores_gemma":[0.086476855,0.00041835415,0.00095033686,0.002907433,0.0016240135,0.0038382972,0.0029531533,0.0015628373,0.0003283832],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032263324,0.0005994994,0.12147575,0.0006037777,0.0011255696,0.00035350857,0.0015650894,0.4557206,0.01992826,0.022376804,0.0038954983,0.36912933],"study_design_scores_gemma":[0.00006259752,0.0003067779,0.012652455,0.000037535003,0.000064335305,0.00019856452,0.0003013264,0.9597103,0.013609391,0.012168536,0.00082614116,0.000062164014],"about_ca_topic_score_codex":0.0026063442,"about_ca_topic_score_gemma":0.0021609508,"teacher_disagreement_score":0.020706441,"about_ca_system_score_codex":0.0018318641,"about_ca_system_score_gemma":0.0012479037,"threshold_uncertainty_score":0.10950744},"labels":[],"label_agreement":null},{"id":"W2538125388","doi":"10.1109/tic-sth.2009.5444360","title":"Direct automatic generation of mind maps from text with M&lt;sup&gt;2&lt;/sup&gt;Gen","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Mind map; Brainstorming; Computer science; Software; Space (punctuation); Concept map; Information retrieval; Topic Maps; World Wide Web; Human–computer interaction; Artificial intelligence; Programming language","score_opus":0.021881827262051776,"score_gpt":0.2227388800389021,"score_spread":0.20085705277685034,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2538125388","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012423742,0.00019777882,0.7170498,0.0003111475,0.00046059914,0.0005760411,0.007687573,0.24108432,0.0202091],"genre_scores_gemma":[0.08063329,0.0002733054,0.84777415,0.00021133345,0.0001680731,0.001631529,0.017310921,0.02661429,0.025383178],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990688,0.000213339,0.00008569389,0.00027526094,0.00029744187,0.000059499213],"domain_scores_gemma":[0.9943786,0.0035997983,0.00021037504,0.00053831516,0.00114115,0.00013162398],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013033784,0.0021263144,0.0007949413,0.003917772,0.0006701263,0.0022764879,0.0012006904,0.0008394427,0.05909829],"category_scores_gemma":[0.010264604,0.00070106913,0.0012192578,0.0018499503,0.00053978,0.0023993477,0.0023590915,0.0007855432,0.028192125],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004671476,0.00015422356,0.0015146064,0.0013152693,0.00011371721,0.00089959893,0.0018358382,0.002535398,0.025867272,0.0075858803,0.14545621,0.81225485],"study_design_scores_gemma":[0.00063563645,0.00037309737,0.007609949,0.0005575461,0.00024701635,0.0021752606,0.002046547,0.25523883,0.17985784,0.050848417,0.5000446,0.00036530904],"about_ca_topic_score_codex":0.0011293978,"about_ca_topic_score_gemma":0.0015104804,"teacher_disagreement_score":0.05909829,"about_ca_system_score_codex":0.00049237965,"about_ca_system_score_gemma":0.0008472832,"threshold_uncertainty_score":0.19770348},"labels":[],"label_agreement":null},{"id":"W2542835211","doi":"10.48550/arxiv.1610.09038","title":"Professor Forcing: A New Algorithm for Training Recurrent Networks","year":2016,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":330,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Forcing (mathematics); Computer science; Treebank; Sampling (signal processing); MNIST database; Algorithm; Handwriting; Sample (material); Artificial intelligence; Sequence (biology); Machine learning; Speech recognition; Artificial neural network; Mathematics; Parsing","score_opus":0.13436034568920915,"score_gpt":0.20817931524737854,"score_spread":0.07381896955816938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2542835211","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066262707,0.00008350059,0.9901815,0.00012906926,0.00003406203,0.000044561133,0.000054731838,0.0018925726,0.00095368444],"genre_scores_gemma":[0.2448703,0.00015145283,0.7476967,0.00027958382,0.000111142326,0.00046332233,0.0005427433,0.00083666155,0.0050481576],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909174,0.0003201575,0.00006232283,0.00022576498,0.0002163354,0.000083786275],"domain_scores_gemma":[0.9973463,0.0015356484,0.00022144041,0.00047217557,0.00033673146,0.00008763255],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025480108,0.0011222567,0.0008595421,0.0009528279,0.000603682,0.0008398616,0.002524542,0.0017036191,0.004368776],"category_scores_gemma":[0.010303588,0.0008306931,0.0009816127,0.00066445774,0.00090042816,0.0019933083,0.001715747,0.0025859566,0.0013959223],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022700673,0.00010472252,0.0017294525,0.00010042822,0.00012664986,0.00015943917,0.00025253417,0.606744,0.010302593,0.02838058,0.006562183,0.34531036],"study_design_scores_gemma":[0.000016729027,0.0000316588,0.000057946974,0.00000642583,0.0000074043423,0.000024660367,0.0000072241346,0.9910966,0.0016040037,0.0061217165,0.0010195887,0.000006036152],"about_ca_topic_score_codex":0.0030824274,"about_ca_topic_score_gemma":0.0066927816,"teacher_disagreement_score":0.004368776,"about_ca_system_score_codex":0.0007792272,"about_ca_system_score_gemma":0.0010770741,"threshold_uncertainty_score":0.014615059},"labels":[],"label_agreement":null},{"id":"W2543408643","doi":"10.1109/iat.2005.50","title":"Category-based Similarity Algorithm for Semantic Similarity in Multi-agent Information Sharing Systems","year":2006,"lang":"en","type":"article","venue":"IEEE/WIC/ACM International Conference on Intelligent Agent Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Semantic similarity; Similarity (geometry); Computer science; Cosine similarity; Information retrieval; Vector space model; Matching (statistics); Similarity measure; Data mining; Artificial intelligence; Pattern recognition (psychology); Mathematics; Image (mathematics)","score_opus":0.08304876257932589,"score_gpt":0.31538465408176813,"score_spread":0.23233589150244224,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2543408643","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0039862,0.00034215383,0.99311197,0.000120707395,0.000066838365,0.0001671242,0.00005004287,0.00030344937,0.0018515043],"genre_scores_gemma":[0.19313322,0.00034647324,0.80270296,0.00013873324,0.000095096686,0.0007040099,0.00035579014,0.00009597955,0.002427704],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952609,0.0016408503,0.00046484283,0.0007329433,0.0016926952,0.0002078343],"domain_scores_gemma":[0.99575245,0.0018245077,0.00026935636,0.00047938764,0.0015281165,0.00014626047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032991108,0.00064339425,0.0015666268,0.003418806,0.001507323,0.0021511272,0.0024710335,0.002035216,0.0029899376],"category_scores_gemma":[0.012744647,0.0003041458,0.0008950961,0.0041108113,0.0012566844,0.004439041,0.0024868194,0.0014993,0.00096369104],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023087866,0.0002482588,0.002024681,0.0004032252,0.0001823155,0.00017909672,0.0005763708,0.19870391,0.003947543,0.29398233,0.0081256665,0.49139574],"study_design_scores_gemma":[0.00003971227,0.000114742936,0.0003741859,0.000035951816,0.00002770455,0.00016324071,0.0001300208,0.84270287,0.0018809338,0.14427716,0.0102184415,0.00003506167],"about_ca_topic_score_codex":0.0036671546,"about_ca_topic_score_gemma":0.0023092676,"teacher_disagreement_score":0.0036671546,"about_ca_system_score_codex":0.00211972,"about_ca_system_score_gemma":0.002022442,"threshold_uncertainty_score":0.01744759},"labels":[],"label_agreement":null},{"id":"W2548734396","doi":"10.5296/jei.v2i2.10040","title":"A Model-Based Method for Content Validation of Automatically Generated Test Items","year":2016,"lang":"en","type":"article","venue":"Journal of Educational Issues","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Test (biology); Computer science; Content (measure theory); Item bank; Item response theory; Information retrieval; Graph; Natural language processing; Artificial intelligence; Data mining; Statistics; Mathematics; Theoretical computer science; Psychometrics","score_opus":0.10287132089019924,"score_gpt":0.3784143698961985,"score_spread":0.27554304900599924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2548734396","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002237638,0.000015375554,0.99458224,0.000031865602,0.0000182234,0.0003128062,0.00014769763,0.0023317446,0.00032239026],"genre_scores_gemma":[0.042703282,0.000024628162,0.9540484,0.000051397878,0.000008971137,0.0013684629,0.0008366762,0.0004842797,0.00047387372],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98513204,0.008282899,0.0008877808,0.0018350695,0.003650847,0.00021133319],"domain_scores_gemma":[0.9471823,0.035085298,0.0018680583,0.0067626303,0.008880486,0.00022127632],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013935901,0.0018629745,0.0009942306,0.0043838974,0.001051443,0.0020334222,0.0026014403,0.0016794492,0.004871882],"category_scores_gemma":[0.084503725,0.0009944201,0.0016854753,0.002204446,0.0010789422,0.0019342969,0.0019245755,0.0030916987,0.0022565543],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004355061,0.0007287497,0.0055683414,0.000697051,0.00051938446,0.00029260744,0.0014974164,0.108578525,0.01811952,0.02776602,0.0090478435,0.826749],"study_design_scores_gemma":[0.00019591051,0.0002498544,0.0025615422,0.0001670352,0.00011408532,0.00023258782,0.00018551142,0.9507983,0.015416987,0.02099822,0.008980702,0.0000993223],"about_ca_topic_score_codex":0.007214903,"about_ca_topic_score_gemma":0.00812469,"teacher_disagreement_score":0.013935901,"about_ca_system_score_codex":0.0019046131,"about_ca_system_score_gemma":0.003051304,"threshold_uncertainty_score":0.073701024},"labels":[],"label_agreement":null},{"id":"W2550969563","doi":"","title":"Online Bayesian Moment Matching for Topic Modeling with Unknown Number of Topics","year":2016,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Latent Dirichlet allocation; Hierarchical Dirichlet process; Hyperparameter; Topic model; Computer science; Dirichlet distribution; Matching (statistics); Dirichlet process; Bayesian probability; Moment (physics); Parametric statistics; Simple (philosophy); Machine learning; Prior probability; Artificial intelligence; Mathematics; Statistics","score_opus":0.024477576970147533,"score_gpt":0.27090619512226605,"score_spread":0.24642861815211853,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2550969563","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003123935,0.000307716,0.9953941,0.00014733896,0.000025170308,0.000030081994,0.000113636364,0.0004514712,0.0004065801],"genre_scores_gemma":[0.19460465,0.0011322983,0.79528725,0.0003550984,0.0004452785,0.000649566,0.0020088246,0.00068406353,0.004832927],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99653924,0.0015703192,0.0001779117,0.0009036454,0.0005925183,0.00021626704],"domain_scores_gemma":[0.9936511,0.0044546975,0.0004735098,0.000830329,0.00043352778,0.00015687235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048880004,0.0010653244,0.002173879,0.0020833644,0.0009993812,0.0016147017,0.0028022365,0.0022661616,0.0046988227],"category_scores_gemma":[0.020576011,0.001313054,0.0016003089,0.0029561203,0.0013807676,0.0052228896,0.0025748645,0.0033422809,0.002330843],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005808195,0.00031796444,0.0019946382,0.0004534352,0.00023099319,0.00029236692,0.0006417193,0.3554786,0.008479263,0.2023569,0.012738475,0.41643482],"study_design_scores_gemma":[0.00002813839,0.000019417032,0.0002743747,0.000016810085,0.000017292132,0.00007051937,0.000023200304,0.8843915,0.0011100927,0.11140684,0.0026133633,0.00002841165],"about_ca_topic_score_codex":0.0035136966,"about_ca_topic_score_gemma":0.0046909014,"teacher_disagreement_score":0.0048880004,"about_ca_system_score_codex":0.0016416915,"about_ca_system_score_gemma":0.0018436717,"threshold_uncertainty_score":0.025850534},"labels":[],"label_agreement":null},{"id":"W2554197048","doi":"","title":"WaterlooClarke: TREC 2015 Contextual Suggestion Track","year":2015,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Task (project management); Track (disk drive); Point of interest; Point (geometry); Contextual design; Information retrieval; Human–computer interaction; Artificial intelligence; Machine learning","score_opus":0.0765092666209366,"score_gpt":0.29243765734721133,"score_spread":0.21592839072627473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2554197048","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040136553,0.01397053,0.10716579,0.016085982,0.004623603,0.0062982757,0.5410218,0.17676096,0.09393649],"genre_scores_gemma":[0.08519561,0.0023464514,0.176748,0.0039339196,0.0007517428,0.0022452206,0.63782805,0.0037105617,0.087240465],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965132,0.0010084471,0.00018467975,0.00060007937,0.0013512769,0.00034227822],"domain_scores_gemma":[0.9899858,0.0020037063,0.00034590985,0.0015812247,0.0053040623,0.00077930617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00598183,0.0015657643,0.0015876445,0.003481928,0.0029320354,0.0027243,0.0022843229,0.002067665,0.03226847],"category_scores_gemma":[0.0128882695,0.0006318673,0.0004598303,0.0034167643,0.0008555455,0.0037852202,0.001990313,0.0024317293,0.018020673],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028109577,0.0001900129,0.0009899621,0.00041955546,0.000046108275,0.000049804796,0.00016299589,0.00075638806,0.004781559,0.0009973402,0.9383735,0.052951667],"study_design_scores_gemma":[0.00039880638,0.00040521476,0.010059191,0.00019416846,0.0001088773,0.00013113361,0.0004644173,0.022019845,0.012793335,0.0022998496,0.95091003,0.00021504749],"about_ca_topic_score_codex":0.28022903,"about_ca_topic_score_gemma":0.54321015,"teacher_disagreement_score":0.28022903,"about_ca_system_score_codex":0.005777915,"about_ca_system_score_gemma":0.009137995,"threshold_uncertainty_score":0.55719584},"labels":[],"label_agreement":null},{"id":"W2557475746","doi":"10.1007/978-3-319-46218-9_2","title":"Argumentation Mining in Parliamentary Discourse","year":2016,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Argumentative; Framing (construction); Argumentation theory; Computer science; Frame (networking); Embedding; Politics; Coding (social sciences); Discourse analysis; Linguistics; Artificial intelligence; Natural language processing; Sociology; Political science; Law; Social science; History","score_opus":0.021645032414913407,"score_gpt":0.2714161143421848,"score_spread":0.2497710819272714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2557475746","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0825551,0.005801009,0.88497275,0.0030199487,0.0003092333,0.00024384873,0.0015285611,0.0020157055,0.01955381],"genre_scores_gemma":[0.63110036,0.0017120631,0.34735644,0.00018559648,0.00038211982,0.0003043656,0.005939985,0.00049722457,0.012521893],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99458706,0.0033778322,0.00037287458,0.0006590182,0.0007640237,0.00023921416],"domain_scores_gemma":[0.97928953,0.01785138,0.00062435796,0.00097070757,0.000970291,0.00029368023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005108075,0.00080061745,0.0012972499,0.0044499002,0.0019008241,0.0046164338,0.0019592629,0.001768468,0.00567946],"category_scores_gemma":[0.029093964,0.0008465919,0.0018260429,0.0041476144,0.0011088321,0.007060542,0.0028790475,0.0028440752,0.0024240066],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006916465,0.00039481258,0.005976458,0.0012946383,0.0003321835,0.0004177623,0.0036591257,0.03585024,0.0043544183,0.14352351,0.01728153,0.78622365],"study_design_scores_gemma":[0.00004741852,0.0000816734,0.0034048278,0.0004142833,0.0001320207,0.0003093542,0.0015058779,0.58126163,0.0056224363,0.3770195,0.03014038,0.000060607406],"about_ca_topic_score_codex":0.0016811467,"about_ca_topic_score_gemma":0.0023637447,"teacher_disagreement_score":0.00567946,"about_ca_system_score_codex":0.0012612982,"about_ca_system_score_gemma":0.0013099537,"threshold_uncertainty_score":0.027014375},"labels":[],"label_agreement":null},{"id":"W2557764419","doi":"10.18653/v1/w17-2623","title":"NewsQA: A Machine Comprehension Dataset","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":742,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Textual entailment; Computer science; Comprehension; Natural language processing; Artificial intelligence; Matching (statistics); Set (abstract data type); Word (group theory); Process (computing); Exploratory analysis; Information retrieval; Machine learning; Logical consequence; Data science; Linguistics","score_opus":0.058483311747084536,"score_gpt":0.30640873117197615,"score_spread":0.2479254194248916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2557764419","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057185207,0.0017579142,0.006196069,0.0016515455,0.00024543572,0.0006567295,0.9136712,0.0069849053,0.011651111],"genre_scores_gemma":[0.032523528,0.0002303554,0.008884204,0.00042316996,0.000102490085,0.00062232366,0.9535761,0.00021354535,0.0034243735],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985902,0.00041219228,0.00018294541,0.0003540167,0.00036237948,0.00009827971],"domain_scores_gemma":[0.99558336,0.0019799348,0.0003145472,0.00061474333,0.0012040704,0.00030347257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013442469,0.001652982,0.00063195586,0.0034376252,0.0008120087,0.0010915934,0.0019190419,0.0023485145,0.012840135],"category_scores_gemma":[0.0084630465,0.00027343738,0.0008662831,0.0024647545,0.00040419048,0.0016974314,0.0011172183,0.0014774104,0.010268539],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039562755,0.00078251754,0.008898137,0.0015397395,0.00014482853,0.0004278564,0.00065519544,0.00272624,0.0053040553,0.0019455077,0.9295832,0.047597077],"study_design_scores_gemma":[0.00087912555,0.00064142706,0.04242887,0.00030693584,0.00018028759,0.0009968383,0.0013504882,0.032841746,0.012838822,0.007178825,0.9001929,0.00016365321],"about_ca_topic_score_codex":0.008136968,"about_ca_topic_score_gemma":0.01808344,"teacher_disagreement_score":0.012840135,"about_ca_system_score_codex":0.0010924247,"about_ca_system_score_gemma":0.0011470835,"threshold_uncertainty_score":0.042954504},"labels":[],"label_agreement":null},{"id":"W2563364981","doi":"10.1007/s10278-016-9931-8","title":"Characterization of Change and Significance for Clinical Findings in Radiology Reports Through Natural Language Processing","year":2017,"lang":"en","type":"article","venue":"Journal of Digital Imaging","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Radiology; Recall; Natural language processing; Artificial intelligence; Precision and recall; Task (project management); Medicine; Health care; Semantics (computer science); Machine learning; Information retrieval; Psychology","score_opus":0.09346756759821408,"score_gpt":0.37056502382546125,"score_spread":0.2770974562272472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2563364981","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8306449,0.005104304,0.14063957,0.0029132734,0.00029981352,0.000708785,0.01357848,0.002421137,0.0036896851],"genre_scores_gemma":[0.9584356,0.000508821,0.03383078,0.00011897746,0.00026999324,0.0001978629,0.006103951,0.00008785146,0.00044621702],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99430436,0.0015764994,0.0013514974,0.0013512281,0.0011092916,0.00030714212],"domain_scores_gemma":[0.9476354,0.038298372,0.0057064896,0.0017311508,0.0057671964,0.00086130213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050306423,0.0005754226,0.000513287,0.012219011,0.0008005316,0.0025468308,0.00093234546,0.001005857,0.001113349],"category_scores_gemma":[0.027491711,0.00029074345,0.0013827323,0.004302118,0.0009986493,0.0029260458,0.001129931,0.0013485613,0.00054669706],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018053417,0.00084303675,0.53534305,0.0017244954,0.0006687057,0.002085438,0.0054560443,0.008866518,0.047502488,0.0063228034,0.01086105,0.37852103],"study_design_scores_gemma":[0.00016030138,0.0010171197,0.60546094,0.00041409806,0.0017983654,0.008141607,0.007202234,0.2933238,0.02938608,0.025461415,0.027353438,0.00028058383],"about_ca_topic_score_codex":0.004673089,"about_ca_topic_score_gemma":0.005367075,"teacher_disagreement_score":0.012219011,"about_ca_system_score_codex":0.0010256349,"about_ca_system_score_gemma":0.0017348255,"threshold_uncertainty_score":0.02660495},"labels":[],"label_agreement":null},{"id":"W2570431255","doi":"10.1162/tacl_a_00077","title":"Aspect-augmented Adversarial Networks for Domain Adaptation","year":2017,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Army Research Office","keywords":"Computer science; Adversarial system; Domain adaptation; Classifier (UML); Artificial intelligence; Transfer of learning; Sentence; Training set; Domain (mathematical analysis); Relevance (law); Natural language processing; Machine learning; Invariant (physics)","score_opus":0.025433414884880583,"score_gpt":0.27190490475990986,"score_spread":0.2464714898750293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2570431255","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009802936,0.0004252828,0.9858505,0.00021484254,0.00007480339,0.000046789602,0.00015027619,0.0009782031,0.0024562373],"genre_scores_gemma":[0.69157255,0.0010182488,0.2935808,0.0007031672,0.0002706828,0.000406338,0.0015004621,0.00039594976,0.0105519],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942625,0.00022169197,0.000026922438,0.00014401072,0.00012984539,0.00005128136],"domain_scores_gemma":[0.9990277,0.00046482962,0.00009732901,0.0002330008,0.00013239181,0.000044715125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011555144,0.0011749084,0.0007708182,0.00054914306,0.00025769812,0.00077452144,0.0013838743,0.00087964896,0.002353759],"category_scores_gemma":[0.0037943241,0.00037351032,0.0007593727,0.0007037246,0.0007265564,0.001420522,0.0019514479,0.002241588,0.0011260015],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000171961,0.000114485745,0.0015722563,0.00013724827,0.00013507971,0.00017643323,0.0001745366,0.73551625,0.0083013475,0.03332645,0.009190292,0.2111837],"study_design_scores_gemma":[0.000005125274,0.000018877745,0.00011461298,0.0000068764575,0.0000075057324,0.00003130525,0.0000065399518,0.9847269,0.0009373505,0.01257219,0.0015664469,0.0000062437875],"about_ca_topic_score_codex":0.0012754264,"about_ca_topic_score_gemma":0.0017589369,"teacher_disagreement_score":0.002353759,"about_ca_system_score_codex":0.0005888284,"about_ca_system_score_gemma":0.000507691,"threshold_uncertainty_score":0.007874131},"labels":[],"label_agreement":null},{"id":"W2572043595","doi":"","title":"WaterlooClarke: TREC 2015 Total Recall Track","year":2015,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Recall; Track (disk drive); Cluster analysis; Selection (genetic algorithm); Process (computing); Precision and recall; Information retrieval; Data mining; Artificial intelligence","score_opus":0.0752741043559129,"score_gpt":0.2872297986682088,"score_spread":0.21195569431229588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2572043595","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047295064,0.035588406,0.0710213,0.04083132,0.020504816,0.007762236,0.47670472,0.046159495,0.2541327],"genre_scores_gemma":[0.093282916,0.0054360544,0.04583829,0.004940753,0.0019363799,0.0023257644,0.5402797,0.0024350034,0.3035253],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99108607,0.0016283192,0.00044088837,0.0007808627,0.0052697654,0.0007940609],"domain_scores_gemma":[0.97139955,0.0022663872,0.000967263,0.0022530563,0.02098079,0.0021328924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014485605,0.002032948,0.002095816,0.0073560355,0.0036163724,0.0069205547,0.0034002734,0.0024039415,0.023699872],"category_scores_gemma":[0.0205852,0.0007930917,0.0009092933,0.0040832716,0.001361721,0.004411845,0.002564009,0.0033809717,0.015864044],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018453778,0.00013713787,0.0007162642,0.0003116088,0.000061445426,0.000033907316,0.000041664687,0.0009562999,0.002010691,0.00095225335,0.9629249,0.031669267],"study_design_scores_gemma":[0.0005675825,0.0007032862,0.021620112,0.0003693745,0.00020898979,0.00025027167,0.00026135036,0.025407037,0.019674994,0.0045824065,0.9260593,0.00029543872],"about_ca_topic_score_codex":0.27196303,"about_ca_topic_score_gemma":0.5217234,"teacher_disagreement_score":0.27196303,"about_ca_system_score_codex":0.012003957,"about_ca_system_score_gemma":0.016343264,"threshold_uncertainty_score":0.54076004},"labels":[],"label_agreement":null},{"id":"W2573266127","doi":"10.63317/2z8qjf6d9e4t","title":"Evaluating a Topic Modelling Approach to Measuring Corpus Similarity","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Similarity (geometry); Natural language processing; Task (project management); Information retrieval; Artificial intelligence; Text corpus; Image (mathematics)","score_opus":0.2631698706386932,"score_gpt":0.31496228741959786,"score_spread":0.05179241678090468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2573266127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46571732,0.006179077,0.51053655,0.00083390577,0.0005771643,0.0010258409,0.0022197524,0.0050127623,0.007897618],"genre_scores_gemma":[0.73766065,0.0012428112,0.2502108,0.00016843561,0.00029990004,0.0006482025,0.0067556165,0.00063290284,0.0023807716],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98277414,0.00901504,0.001477639,0.0024355445,0.0038594776,0.0004381725],"domain_scores_gemma":[0.93048984,0.05590816,0.0015217129,0.0035396947,0.0073779942,0.0011625228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01596065,0.0013648318,0.0016285992,0.009953908,0.0017805208,0.005506328,0.0021350333,0.0032527107,0.0018650522],"category_scores_gemma":[0.081373036,0.00053689216,0.0018987292,0.007547434,0.00083386933,0.006312611,0.0032187216,0.001904876,0.001233002],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0041668755,0.0014275123,0.054565072,0.001727426,0.0024564718,0.00044378138,0.003459578,0.08280246,0.027363144,0.00831169,0.012743095,0.8005329],"study_design_scores_gemma":[0.00022883278,0.0012950165,0.021606375,0.00013044082,0.0009890913,0.0006800659,0.0016173981,0.94157535,0.015784765,0.009619833,0.0063311234,0.00014164351],"about_ca_topic_score_codex":0.008200644,"about_ca_topic_score_gemma":0.0091925515,"teacher_disagreement_score":0.01596065,"about_ca_system_score_codex":0.0020551065,"about_ca_system_score_gemma":0.002291714,"threshold_uncertainty_score":0.084409},"labels":[],"label_agreement":null},{"id":"W2573346038","doi":"","title":"Predicting sentential semantic compatibility for aggregation in text-to-text generation","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Sentence; Natural language processing; Artificial intelligence; Task (project management); Cluster analysis; Process (computing); Context (archaeology); Information retrieval","score_opus":0.07688502906661919,"score_gpt":0.3287613563504564,"score_spread":0.2518763272838372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2573346038","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63778776,0.001482883,0.33887175,0.0009890356,0.00028972852,0.00074766664,0.003938065,0.008942451,0.0069507193],"genre_scores_gemma":[0.81068647,0.00018119685,0.18186313,0.00007538965,0.000095577416,0.00021795338,0.0057386192,0.00033899726,0.0008026836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99631566,0.0016773539,0.00033832705,0.00085574825,0.00059032295,0.0002225627],"domain_scores_gemma":[0.9759448,0.017852766,0.0013669856,0.0016041829,0.0027178419,0.00051334756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056519513,0.0011222285,0.0007154646,0.003934626,0.0011307751,0.0019804933,0.0009407126,0.0013963438,0.001993907],"category_scores_gemma":[0.034592204,0.00037146654,0.0009722483,0.0023721163,0.00051181694,0.003942818,0.0014040117,0.0013635196,0.0012749054],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022870337,0.0010424894,0.077985264,0.0011454892,0.00039425585,0.000852232,0.0029683989,0.06360371,0.057282366,0.010073823,0.027827581,0.7545374],"study_design_scores_gemma":[0.00010224182,0.00036830318,0.036866017,0.00006849034,0.00019349463,0.00037216058,0.0008095256,0.89315146,0.044634435,0.01722133,0.006120567,0.00009199715],"about_ca_topic_score_codex":0.003557357,"about_ca_topic_score_gemma":0.0060920073,"teacher_disagreement_score":0.0056519513,"about_ca_system_score_codex":0.00091341796,"about_ca_system_score_gemma":0.001288601,"threshold_uncertainty_score":0.029890716},"labels":[],"label_agreement":null},{"id":"W2573425638","doi":"","title":"Distraction-based neural networks for modeling documents","year":2016,"lang":"en","type":"article","venue":"International Joint Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":79,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Automatic summarization; Computer science; Distraction; GRASP; Artificial neural network; Representation (politics); Artificial intelligence; Natural language processing; Traverse; Information retrieval; Machine learning; Programming language","score_opus":0.1398677867114096,"score_gpt":0.34046305885792866,"score_spread":0.20059527214651907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2573425638","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12772131,0.0031662274,0.8617611,0.0010597501,0.00023233042,0.00013941083,0.0008907966,0.0018947797,0.0031344066],"genre_scores_gemma":[0.862803,0.0012490064,0.12544505,0.00034613136,0.00022665909,0.00041089382,0.0018426032,0.000112878304,0.007563742],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996195,0.00013487264,0.000024256535,0.0001268029,0.000053709955,0.000040711133],"domain_scores_gemma":[0.99886847,0.0007273203,0.00012297425,0.0000871127,0.00015579152,0.000038368154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011335253,0.0010051597,0.0007083417,0.001234916,0.00036108715,0.00085485214,0.001705898,0.0013670373,0.0017060611],"category_scores_gemma":[0.004334306,0.00038357527,0.0007674712,0.0016251914,0.00048569418,0.0018009157,0.0006992488,0.0022247636,0.00056633784],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017885055,0.000106439955,0.0016602885,0.000104762876,0.00006743332,0.0000801744,0.00014181367,0.9070443,0.0023107973,0.010588394,0.0021603347,0.07555647],"study_design_scores_gemma":[0.000004975131,0.00001118073,0.00010162774,0.0000040220507,0.00000500141,0.00000647212,0.000004694368,0.99548304,0.00025474982,0.0038269819,0.00029405532,0.0000032501248],"about_ca_topic_score_codex":0.0075603854,"about_ca_topic_score_gemma":0.009002744,"teacher_disagreement_score":0.0075603854,"about_ca_system_score_codex":0.0015077385,"about_ca_system_score_gemma":0.000487492,"threshold_uncertainty_score":0.015032709},"labels":[],"label_agreement":null},{"id":"W2574914175","doi":"","title":"An interactive system for exploring community question answering forums","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Question answering; Computer science; World Wide Web; Interface (matter); Information retrieval; User interface; Graphical user interface; Human–computer interaction","score_opus":0.13824046408852536,"score_gpt":0.3595167125248636,"score_spread":0.22127624843633825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2574914175","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035083596,0.00080519036,0.6966871,0.0005210279,0.00020263718,0.0011284478,0.009824833,0.24069145,0.015055738],"genre_scores_gemma":[0.20685954,0.00046488404,0.73389816,0.00040864918,0.0003004574,0.0028826725,0.03027757,0.008683364,0.016224628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985973,0.00049616955,0.0000888418,0.00033512723,0.00036514679,0.00011751255],"domain_scores_gemma":[0.9946569,0.0034300706,0.00016950561,0.0005095545,0.0006763539,0.00055762863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002468556,0.0013142282,0.0008115504,0.0034324701,0.0013826342,0.0019083596,0.0021866944,0.001540647,0.03620894],"category_scores_gemma":[0.008234875,0.0006387189,0.0007924497,0.0019283043,0.00044117324,0.004337752,0.004600656,0.00094814505,0.008040996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002829979,0.0008219493,0.005865743,0.0019018976,0.0002307103,0.0012350305,0.00804736,0.004549758,0.068555154,0.016047858,0.19188261,0.69803196],"study_design_scores_gemma":[0.0013456183,0.0012022217,0.009638632,0.00041425615,0.00034152757,0.0020825847,0.0027726584,0.22402844,0.052937668,0.04358964,0.66110545,0.00054123456],"about_ca_topic_score_codex":0.001617231,"about_ca_topic_score_gemma":0.0022260745,"teacher_disagreement_score":0.03620894,"about_ca_system_score_codex":0.0005166664,"about_ca_system_score_gemma":0.00073869183,"threshold_uncertainty_score":0.12113094},"labels":[],"label_agreement":null},{"id":"W2577720462","doi":"","title":"Named Entity Disambiguation for little known referents: a topic-based approach","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Referent; Exploit; Task (project management); Property (philosophy); Named-entity recognition; Natural language processing; Entity linking; Artificial intelligence; Information retrieval; Training set; Named entity; Linguistics; Knowledge base","score_opus":0.09194997485170264,"score_gpt":0.3341074863322806,"score_spread":0.24215751148057796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2577720462","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011312237,0.0013296204,0.9804185,0.0007875542,0.00022014102,0.00012867694,0.00072005665,0.0032065401,0.0018767164],"genre_scores_gemma":[0.24928297,0.0016733521,0.7327743,0.00069201266,0.001032024,0.0004247628,0.0061770305,0.0010457304,0.0068977415],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941949,0.0017809536,0.0005309291,0.0020761576,0.0011027304,0.00031428356],"domain_scores_gemma":[0.98966837,0.0056091757,0.00068763475,0.0020181723,0.0017111066,0.0003054929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062201777,0.0016567499,0.0024793586,0.009488828,0.002800216,0.0037779182,0.0046927542,0.0029983376,0.0030813725],"category_scores_gemma":[0.014429288,0.0012298749,0.0027074856,0.008121777,0.0012881155,0.009456827,0.0049225083,0.003011609,0.0044672983],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062236085,0.00066686433,0.0112306615,0.0011153544,0.00065530796,0.0013828077,0.004280854,0.03827206,0.030209402,0.0416601,0.033995934,0.8359083],"study_design_scores_gemma":[0.000114797476,0.00019709943,0.00552349,0.00029066138,0.0006398302,0.001979173,0.0020648604,0.7761338,0.032353573,0.096247174,0.0842006,0.00025507508],"about_ca_topic_score_codex":0.0032183377,"about_ca_topic_score_gemma":0.006890372,"teacher_disagreement_score":0.009488828,"about_ca_system_score_codex":0.00091281027,"about_ca_system_score_gemma":0.002376503,"threshold_uncertainty_score":0.032895803},"labels":[],"label_agreement":null},{"id":"W2578302695","doi":"10.1109/wi.2016.0095","title":"Context Free Frequently Asked Questions Detection Using Machine Learning Techniques","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Fredericton","funders":"","keywords":"Computer science; Frequently asked questions; Naive Bayes classifier; Classifier (UML); Artificial intelligence; Information retrieval; Machine learning; Natural language processing; Support vector machine; Feature engineering; Parsing","score_opus":0.027526159804951964,"score_gpt":0.2582787823248984,"score_spread":0.23075262251994647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2578302695","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23390287,0.0039207255,0.74417067,0.00075522193,0.0003018021,0.00087915105,0.003914622,0.0088933315,0.0032615927],"genre_scores_gemma":[0.70911187,0.00065083726,0.2803584,0.00024105328,0.0003233028,0.0005748312,0.0060251993,0.00012820581,0.0025863177],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959645,0.001296716,0.00045210661,0.0010477673,0.00091170275,0.00032717822],"domain_scores_gemma":[0.99164283,0.0048101423,0.0010400778,0.00056446163,0.0016794387,0.0002630411],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026893073,0.0011947476,0.0010574638,0.007084907,0.0008537763,0.0014336667,0.0012370069,0.0018324945,0.0014650277],"category_scores_gemma":[0.009217928,0.00030678048,0.0014750059,0.0030033248,0.00037643022,0.0022065928,0.0010814323,0.0012610051,0.001306381],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005975601,0.00084820937,0.05127933,0.0008044689,0.000323112,0.0013499385,0.001555859,0.015956016,0.061822586,0.004369969,0.016213035,0.84488004],"study_design_scores_gemma":[0.000061551975,0.00051121024,0.07016423,0.00014081757,0.0002006304,0.0022251506,0.001499362,0.8500414,0.0387424,0.014415003,0.021817915,0.00018031723],"about_ca_topic_score_codex":0.0032621718,"about_ca_topic_score_gemma":0.0038522205,"teacher_disagreement_score":0.007084907,"about_ca_system_score_codex":0.0006690192,"about_ca_system_score_gemma":0.00080197124,"threshold_uncertainty_score":0.014222562},"labels":[],"label_agreement":null},{"id":"W2578945560","doi":"","title":"Real Time Filtering of Tweets Using Wikipedia Concepts and Google Tri-gram Semantic Relatedness","year":2015,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Information retrieval; n-gram; Semantic similarity; Word (group theory); Social media; World Wide Web; Natural language processing; Language model","score_opus":0.07837344284233229,"score_gpt":0.3094833050650553,"score_spread":0.231109862222723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2578945560","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5646859,0.003639693,0.3740354,0.001137115,0.0009863785,0.001555434,0.01720816,0.023228139,0.013523793],"genre_scores_gemma":[0.6769998,0.0007001322,0.29394427,0.00016042205,0.00034456604,0.0005274453,0.023347905,0.00042399607,0.0035514778],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986695,0.00026598995,0.00013766621,0.00036812975,0.0004247991,0.0001339713],"domain_scores_gemma":[0.9982017,0.00056246243,0.00020369564,0.0001599977,0.0007584417,0.00011357433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010199769,0.0014644064,0.0010455328,0.011258868,0.001144069,0.0014291485,0.00064288307,0.00089624006,0.0016948178],"category_scores_gemma":[0.0051731835,0.00026930476,0.0010270072,0.0051656477,0.00030107467,0.002671326,0.0009942116,0.0007721198,0.002620385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023627854,0.0014235604,0.027475985,0.0017098903,0.0005965793,0.0014304534,0.0020940176,0.016955722,0.10253004,0.0054495474,0.0458358,0.7921356],"study_design_scores_gemma":[0.00013487572,0.0009880048,0.046345927,0.00018294965,0.0003468475,0.0011772952,0.0024557095,0.84491026,0.053524908,0.012338354,0.037339747,0.00025504464],"about_ca_topic_score_codex":0.010838072,"about_ca_topic_score_gemma":0.016243283,"teacher_disagreement_score":0.011258868,"about_ca_system_score_codex":0.0007016975,"about_ca_system_score_gemma":0.0009506127,"threshold_uncertainty_score":0.02154994},"labels":[],"label_agreement":null},{"id":"W2579773546","doi":"","title":"Reddit Temporal N-gram Corpus and its Applications on Paraphrase and Semantic Similarity in Social Media using a Topic-based Latent Semantic Analysis.","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Paraphrase; Computer science; Latent semantic analysis; SemEval; Natural language processing; Artificial intelligence; Social media; Semantic similarity; Similarity (geometry); n-gram; Probabilistic latent semantic analysis; Text corpus; Information retrieval; Language model; World Wide Web","score_opus":0.08700156074749933,"score_gpt":0.3337642859727162,"score_spread":0.24676272522521686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2579773546","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44182605,0.004779739,0.40748218,0.0024836445,0.001251,0.0017842584,0.09786966,0.016398208,0.026125243],"genre_scores_gemma":[0.5147017,0.0009896682,0.33904642,0.00026228226,0.0003332823,0.002542671,0.13378134,0.0009644694,0.0073781135],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981389,0.00089813635,0.00015492948,0.00030378852,0.00042507253,0.000079271784],"domain_scores_gemma":[0.9943094,0.0029715463,0.00032998182,0.0010660248,0.0011261891,0.00019683663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014649419,0.000792904,0.00053502957,0.004571416,0.0017057266,0.0010965274,0.0011068844,0.00096268795,0.005859961],"category_scores_gemma":[0.012498006,0.00028789026,0.0005627694,0.004906395,0.00072823954,0.002960093,0.002502844,0.001221627,0.003157371],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021067155,0.0013249174,0.01241132,0.0031779355,0.00025079542,0.002362316,0.004167434,0.016675029,0.061229795,0.0300867,0.14024632,0.7259607],"study_design_scores_gemma":[0.00038205352,0.0007674061,0.044630487,0.00048548458,0.00016786908,0.0033904633,0.0043852073,0.6219219,0.06434038,0.042790443,0.21640138,0.00033693798],"about_ca_topic_score_codex":0.0043808077,"about_ca_topic_score_gemma":0.009902752,"teacher_disagreement_score":0.005859961,"about_ca_system_score_codex":0.0007703181,"about_ca_system_score_gemma":0.0011198667,"threshold_uncertainty_score":0.01960355},"labels":[],"label_agreement":null},{"id":"W2583976214","doi":"10.1145/3018661.3018692","title":"Document Retrieval Model Through Semantic Linking","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Information retrieval; Explicit semantic analysis; Complement (music); Graph; Semantic similarity; Probabilistic logic; Natural language processing; Semantic computing; Artificial intelligence; Semantic technology; Theoretical computer science; Semantic Web","score_opus":0.05264179382538973,"score_gpt":0.30707330122771526,"score_spread":0.2544315074023255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2583976214","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012323813,0.0006679345,0.97891825,0.00060772523,0.000063948115,0.00012212005,0.0003233165,0.0009861556,0.0059866207],"genre_scores_gemma":[0.4625645,0.0023878398,0.5111164,0.00046475,0.0004289128,0.0007776465,0.0021052277,0.00041059774,0.019744162],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981365,0.00064487127,0.00014007442,0.00043290525,0.00055131444,0.000094393916],"domain_scores_gemma":[0.9982389,0.0010140317,0.00017397194,0.00025670274,0.00026266446,0.000053751453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020333477,0.00069535733,0.0008698749,0.0034015826,0.0007140035,0.0026764418,0.0020840336,0.0014307942,0.0033756755],"category_scores_gemma":[0.006116629,0.00038328007,0.001620056,0.0033521994,0.00094735593,0.0074001797,0.0014669012,0.0010743497,0.0023373899],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002584398,0.00023225616,0.0013164283,0.00048851484,0.00017406617,0.00041640818,0.0010937444,0.3277193,0.008962981,0.4647858,0.008302249,0.18624988],"study_design_scores_gemma":[0.000043481487,0.00008670713,0.00026429218,0.000029537914,0.000079916834,0.0002936097,0.00006265527,0.8651156,0.002155591,0.119334824,0.012485274,0.000048571634],"about_ca_topic_score_codex":0.002381744,"about_ca_topic_score_gemma":0.0014132041,"teacher_disagreement_score":0.0034015826,"about_ca_system_score_codex":0.0013181871,"about_ca_system_score_gemma":0.0010838448,"threshold_uncertainty_score":0.011292756},"labels":[],"label_agreement":null},{"id":"W2584220694","doi":"10.1609/aaai.v32i1.11321","title":"RUBER: An Unsupervised Method for Automatic Evaluation of Open-Domain Dialog Systems","year":2018,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":162,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China; Tencent; National Science Foundation","keywords":"Dialog box; Computer science; Annotation; Metric (unit); Open domain; Artificial intelligence; Utterance; Domain (mathematical analysis); Natural language processing; Dialog system; Transferability; Conversation; Information retrieval; Machine learning; World Wide Web; Linguistics; Mathematics","score_opus":0.24357062726124815,"score_gpt":0.40746944321116096,"score_spread":0.16389881594991282,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2584220694","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018759256,0.000873649,0.9121875,0.00015596156,0.00021258116,0.00086179655,0.00425271,0.058453888,0.004242569],"genre_scores_gemma":[0.21574593,0.00025985835,0.75768465,0.0002211047,0.00015071488,0.0020291654,0.01354065,0.0042313305,0.006136598],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9802033,0.009049145,0.0017738988,0.0044495347,0.0038574808,0.00066671986],"domain_scores_gemma":[0.98418945,0.0063555124,0.0011651951,0.0032336512,0.004440743,0.0006154997],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010365751,0.0034801206,0.0016399496,0.0068818056,0.0012934154,0.0028776391,0.0035780305,0.0022827848,0.0065282215],"category_scores_gemma":[0.03332132,0.00072465965,0.0014594357,0.0026709717,0.0010262633,0.004353632,0.0047396366,0.002548281,0.0073789777],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077920005,0.0004602764,0.005079574,0.0014278381,0.00045655298,0.00020450017,0.0017469849,0.015894953,0.03504214,0.0093045635,0.048210744,0.88139266],"study_design_scores_gemma":[0.00020740632,0.000726643,0.011364924,0.00024600836,0.00016952389,0.00062392076,0.0011800678,0.8185634,0.069659404,0.032227036,0.06461286,0.00041883206],"about_ca_topic_score_codex":0.0030475378,"about_ca_topic_score_gemma":0.0056651896,"teacher_disagreement_score":0.010365751,"about_ca_system_score_codex":0.0013847525,"about_ca_system_score_gemma":0.0023178593,"threshold_uncertainty_score":0.05482},"labels":[],"label_agreement":null},{"id":"W2584524946","doi":"10.1109/bigdata.2016.7841060","title":"Efficient natural language pre-processing for analyzing large data sets","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Artificial intelligence; WordNet; Pipeline (software); Natural language processing; Machine translation; Preprocessor; Task (project management); Graph; Variety (cybernetics); Machine learning","score_opus":0.02934926712142034,"score_gpt":0.3177701427171664,"score_spread":0.2884208755957461,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2584524946","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0069711837,0.0006150529,0.97485036,0.00049958547,0.00012417117,0.0008209222,0.00296671,0.011798427,0.0013536009],"genre_scores_gemma":[0.031136386,0.00069568615,0.9557526,0.00022516168,0.00018732257,0.0015355938,0.008723563,0.00050308043,0.0012405798],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99729306,0.0007293119,0.0003460462,0.00063201913,0.00086218544,0.00013733075],"domain_scores_gemma":[0.98890734,0.006757912,0.000828155,0.0016317056,0.0017189195,0.00015601641],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002873845,0.0024433879,0.0017630744,0.0061291885,0.0016731643,0.0028713702,0.001849284,0.001065364,0.0066622472],"category_scores_gemma":[0.011691533,0.0010859831,0.0026224146,0.005050937,0.0013739204,0.0038678327,0.0024344039,0.0030240086,0.011280659],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002733251,0.00033550378,0.0024108083,0.0018777363,0.00017281657,0.00076442683,0.0007767529,0.011014015,0.09251669,0.011628565,0.022108003,0.85612124],"study_design_scores_gemma":[0.00012253644,0.0006662197,0.015040083,0.00045042904,0.00030676543,0.0020554238,0.0026980625,0.59859085,0.12368268,0.13622248,0.11985641,0.00030802574],"about_ca_topic_score_codex":0.0030014107,"about_ca_topic_score_gemma":0.004739769,"teacher_disagreement_score":0.0066622472,"about_ca_system_score_codex":0.0010525134,"about_ca_system_score_gemma":0.0026703163,"threshold_uncertainty_score":0.022287428},"labels":[],"label_agreement":null},{"id":"W2589049976","doi":"10.3390/data2010011","title":"Erratum: Morrison, H., et al. Open Access Article Processing Charges (OA APC) Longitudinal Study 2015 Preliminary Dataset","year":2017,"lang":"en","type":"erratum","venue":"Data","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Data science","score_opus":0.27674736986145604,"score_gpt":0.4640049221509289,"score_spread":0.18725755228947288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2589049976","genre_codex":"editorial","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012721707,0.0015336858,0.0024479695,0.16591965,0.62606204,0.000251693,0.19292758,0.001838467,0.007746797],"genre_scores_gemma":[0.033523668,0.007497468,0.015876196,0.2928137,0.10041962,0.0031009929,0.35429537,0.0069670076,0.18550597],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9791241,0.003619606,0.006435886,0.0016879236,0.008196897,0.00093556027],"domain_scores_gemma":[0.82226795,0.066800945,0.0118670985,0.012178058,0.08410041,0.0027855772],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.01871449,0.001797266,0.0018360364,0.0069599026,0.003627134,0.0048413393,0.003322264,0.003170786,0.061427284],"category_scores_gemma":[0.23507884,0.0016634159,0.002074171,0.009176154,0.001477848,0.002993875,0.003291044,0.008185815,0.044353243],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022079103,0.0000034574468,0.00020168899,0.00005640387,0.000007713146,0.00001673524,0.000010539586,0.000014652887,0.000005407931,0.00012682173,0.9983531,0.001181308],"study_design_scores_gemma":[0.00015431145,0.000025873216,0.004375992,0.0010430184,0.00010555472,0.00018056552,0.00023877945,0.00023752484,0.00028149344,0.0016283095,0.99164796,0.000080484184],"about_ca_topic_score_codex":0.053144176,"about_ca_topic_score_gemma":0.06876518,"teacher_disagreement_score":0.9951587,"about_ca_system_score_codex":0.006347722,"about_ca_system_score_gemma":0.013277466,"threshold_uncertainty_score":0.2054947},"labels":[],"label_agreement":null},{"id":"W2592420180","doi":"10.5220/0006256507380745","title":"Dynamic Selection of Exemplar-SVMs for Watch-list Screening through Domain Adaptation","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université du Québec","funders":"","keywords":"Domain adaptation; Support vector machine; Computer science; Selection (genetic algorithm); Artificial intelligence; Adaptation (eye); Machine learning; Domain (mathematical analysis); Pattern recognition (psychology); Mathematics; Classifier (UML); Psychology","score_opus":0.05166690286437465,"score_gpt":0.3054666274369633,"score_spread":0.25379972457258865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2592420180","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20303908,0.005003477,0.76847225,0.001400202,0.0008932776,0.00046001465,0.0022090415,0.013941599,0.004580941],"genre_scores_gemma":[0.78878576,0.0008929035,0.19100328,0.00068902294,0.000586394,0.0004039293,0.010051215,0.0008785348,0.0067090304],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981996,0.0005047749,0.00013482421,0.00059374643,0.00027919738,0.00028782693],"domain_scores_gemma":[0.99631643,0.0016590273,0.00014563603,0.00040719673,0.001137274,0.00033453337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022929355,0.0018906209,0.0024566504,0.0029233315,0.0009510267,0.001759751,0.003009293,0.0024648788,0.0041795163],"category_scores_gemma":[0.008859234,0.0005177976,0.001800637,0.001826794,0.00047674784,0.0025908058,0.0019184969,0.0028308884,0.0043039015],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015972949,0.0013586779,0.014858412,0.00035954954,0.0003920683,0.00029375232,0.00025364605,0.06388982,0.025587793,0.0023854615,0.047635105,0.84138834],"study_design_scores_gemma":[0.000055324006,0.00012664298,0.001841046,0.00003108099,0.00011588277,0.00011434324,0.000109827415,0.9879743,0.0053991554,0.0021332458,0.0020743238,0.000024826853],"about_ca_topic_score_codex":0.004955481,"about_ca_topic_score_gemma":0.007033766,"teacher_disagreement_score":0.004955481,"about_ca_system_score_codex":0.0006318363,"about_ca_system_score_gemma":0.0015739519,"threshold_uncertainty_score":0.013981819},"labels":[],"label_agreement":null},{"id":"W2593751037","doi":"10.5087/dad.2017.102","title":"Training End-to-End Dialogue Systems with the Ubuntu Dialogue Corpus","year":2017,"lang":"en","type":"article","venue":"Dialogue & Discourse","topic":"Topic Modeling","field":"Computer Science","cited_by":140,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Samsung; Natural Sciences and Engineering Research Council of Canada; Samsung Advanced Institute of Technology","keywords":"Utterance; Computer science; Conversation; Context (archaeology); Task (project management); Artificial intelligence; Construct (python library); Natural language processing; Feature engineering; End-to-end principle; Recall; Feature (linguistics); Precision and recall; Artificial neural network; Speech recognition; Machine learning; Deep learning; Linguistics","score_opus":0.049103931591809494,"score_gpt":0.28453779010342706,"score_spread":0.23543385851161758,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2593751037","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38809693,0.004139055,0.47180998,0.0013035857,0.0020741478,0.002080805,0.012063,0.099728346,0.018704139],"genre_scores_gemma":[0.59135574,0.0005424359,0.35539946,0.0008249721,0.00018913239,0.0022619504,0.032201547,0.0022631232,0.014961687],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99726546,0.0012059776,0.00012510442,0.0009867853,0.00022822796,0.00018849515],"domain_scores_gemma":[0.9962167,0.0021602383,0.000109800996,0.0005528635,0.0007625284,0.00019782412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034141045,0.0027858817,0.0011831959,0.00091173203,0.0010308787,0.0016915674,0.0026441691,0.002213203,0.0067166337],"category_scores_gemma":[0.010765325,0.0008683479,0.0011433831,0.00062724965,0.000728379,0.0028759227,0.0021832222,0.0034330576,0.005764708],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025516942,0.002400896,0.0063761603,0.0014886844,0.00072496384,0.0007614478,0.0024470075,0.33147994,0.023586055,0.003604995,0.07930416,0.5452741],"study_design_scores_gemma":[0.0001818157,0.0004252322,0.0020235768,0.000089312714,0.00009030511,0.0001504169,0.0004735979,0.96206075,0.01693306,0.0029817547,0.014509575,0.000080688005],"about_ca_topic_score_codex":0.011462662,"about_ca_topic_score_gemma":0.016107159,"teacher_disagreement_score":0.011462662,"about_ca_system_score_codex":0.001702836,"about_ca_system_score_gemma":0.0012767372,"threshold_uncertainty_score":0.022791862},"labels":[],"label_agreement":null},{"id":"W2594165071","doi":"10.1007/978-3-319-59569-6_50","title":"Vector Embedding of Wikipedia Concepts and Entities","year":2017,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Embedding; Word embedding; Analogy; Similarity (geometry); Word (group theory); Natural language processing; Artificial intelligence; Task (project management); Vector space; Information retrieval; Linguistics; Image (mathematics); Mathematics","score_opus":0.023726935634417114,"score_gpt":0.2810208361709135,"score_spread":0.25729390053649637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2594165071","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08094344,0.015065317,0.84746385,0.0018509241,0.0020755536,0.00020204602,0.022868976,0.003932975,0.025596889],"genre_scores_gemma":[0.54115725,0.012969799,0.3550636,0.00029719027,0.0009285917,0.00037133484,0.049691055,0.0008961601,0.038625028],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994929,0.00012737632,0.000046196572,0.00016361901,0.00013804606,0.000031741143],"domain_scores_gemma":[0.99911195,0.00030578766,0.00009127606,0.00013156366,0.00030975384,0.00004964335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004901257,0.00069318316,0.0004482831,0.0033727041,0.00031682875,0.0015562612,0.00054972485,0.00049728126,0.004083605],"category_scores_gemma":[0.003036631,0.0002678664,0.0004878155,0.0047289142,0.00030465427,0.0033595962,0.00094572135,0.00080005697,0.0024759383],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003202136,0.00019143161,0.0037657684,0.0014930664,0.00022257383,0.000232806,0.00073617353,0.015202244,0.0175672,0.067797974,0.09121461,0.801256],"study_design_scores_gemma":[0.000051329138,0.00020497614,0.012868843,0.0008637599,0.00031451546,0.0017112856,0.000910125,0.50212944,0.019880155,0.18067558,0.28023502,0.00015504597],"about_ca_topic_score_codex":0.001867406,"about_ca_topic_score_gemma":0.0032665874,"teacher_disagreement_score":0.004083605,"about_ca_system_score_codex":0.00043248365,"about_ca_system_score_gemma":0.00054130505,"threshold_uncertainty_score":0.013661027},"labels":[],"label_agreement":null},{"id":"W2594363579","doi":"10.1016/j.ipm.2017.02.009","title":"Self-training on refined clause patterns for relation extraction","year":2017,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Relationship extraction; Computer science; Parsing; Information extraction; Natural language processing; Dependency (UML); Context (archaeology); Relation (database); Dependency grammar; Benchmark (surveying); Variety (cybernetics); Artificial intelligence; Rule-based machine translation; Data mining","score_opus":0.045633253515832285,"score_gpt":0.3055209064605257,"score_spread":0.2598876529446934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2594363579","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.090051174,0.0012078239,0.86574453,0.00079601206,0.0004213075,0.00057424343,0.0039351867,0.029389653,0.00788014],"genre_scores_gemma":[0.35902825,0.00049982767,0.6088818,0.00054966524,0.00020162552,0.0004225118,0.018422669,0.0016042453,0.010389479],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99851865,0.00032808454,0.00016425132,0.00056273676,0.00025560876,0.00017073915],"domain_scores_gemma":[0.9946896,0.003081423,0.00018897753,0.0010467146,0.00084045343,0.00015276103],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018216105,0.0013297728,0.00092571246,0.002613437,0.0006984726,0.0011381865,0.002101846,0.0014751077,0.015719555],"category_scores_gemma":[0.007977788,0.00076494046,0.001298217,0.0024392875,0.00043514872,0.0039982283,0.0020875412,0.002437027,0.007849737],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042989512,0.00035240303,0.005010174,0.00033379838,0.00013338782,0.00023570197,0.000412582,0.007697076,0.020696959,0.003558997,0.028331183,0.9328078],"study_design_scores_gemma":[0.0001731538,0.0003292686,0.0052366047,0.00019716691,0.0003101108,0.00076600077,0.00070391723,0.8986702,0.03998945,0.021562457,0.032009214,0.000052438656],"about_ca_topic_score_codex":0.004323009,"about_ca_topic_score_gemma":0.009098784,"teacher_disagreement_score":0.015719555,"about_ca_system_score_codex":0.00045686608,"about_ca_system_score_gemma":0.0014509859,"threshold_uncertainty_score":0.05258715},"labels":[],"label_agreement":null},{"id":"W2596863005","doi":"10.1007/978-3-319-51049-1","title":"Prediction and Inference from Social Networks and Social Media","year":2017,"lang":"en","type":"book","venue":"Lecture notes in social networks","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Inference; Social media; Social network analysis; Computer science; Data science; Social network (sociolinguistics); Artificial intelligence; Sociology; World Wide Web","score_opus":0.0274515790927396,"score_gpt":0.25335123633199763,"score_spread":0.22589965723925803,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2596863005","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02222504,0.0035082474,0.9657012,0.0014675156,0.0002658989,0.000077006494,0.001354089,0.0018340935,0.0035669524],"genre_scores_gemma":[0.52217865,0.008909272,0.43462723,0.00034329423,0.0019765848,0.00043716846,0.010243894,0.00051303266,0.020770919],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9991009,0.00028944036,0.000052021736,0.00028155206,0.00021871473,0.000057344205],"domain_scores_gemma":[0.99495065,0.004237885,0.00015573672,0.00035516726,0.00021373833,0.00008689159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018204916,0.001425203,0.0012270795,0.0022436045,0.0005149269,0.0017186169,0.0015609085,0.0013203897,0.0028700926],"category_scores_gemma":[0.0085134925,0.0010084142,0.0016618598,0.0023138258,0.0005641768,0.0041003167,0.0014241704,0.0025207,0.0018786784],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002415893,0.00034000902,0.0069777463,0.000428892,0.00039672802,0.00026570566,0.0002747478,0.3074455,0.0028000511,0.030897208,0.03360894,0.6163228],"study_design_scores_gemma":[0.0000069363264,0.000015833166,0.0006217653,0.00001878905,0.000033195734,0.000041361323,0.000023744562,0.946097,0.0006433989,0.050477475,0.0020113687,0.000009104404],"about_ca_topic_score_codex":0.004885466,"about_ca_topic_score_gemma":0.0066511645,"teacher_disagreement_score":0.004885466,"about_ca_system_score_codex":0.0007176576,"about_ca_system_score_gemma":0.0005705579,"threshold_uncertainty_score":0.009714067},"labels":[],"label_agreement":null},{"id":"W2602143361","doi":"10.1109/icsc.2017.9","title":"A Context-Aware Approach for the Identification of Complex Words in Natural Language Texts","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Identification (biology); Natural language processing; Artificial intelligence; Context (archaeology); SemEval; Word (group theory); Natural language; Natural language understanding; Natural (archaeology); Word identification; Linguistics; Word recognition; Task (project management)","score_opus":0.056285165004476466,"score_gpt":0.31627142238983447,"score_spread":0.259986257385358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2602143361","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35027453,0.017996823,0.6033819,0.0007343775,0.0007128302,0.0009297716,0.0037357917,0.013300474,0.008933594],"genre_scores_gemma":[0.6862484,0.0017592451,0.30299532,0.00025309896,0.00059787434,0.0003840998,0.0046737543,0.00036613087,0.0027221343],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99780077,0.0006743545,0.0001884719,0.00079521484,0.00040763113,0.00013343123],"domain_scores_gemma":[0.99434596,0.0032669576,0.0005470866,0.0005807398,0.0010592587,0.00020009668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015040225,0.001420332,0.0008894666,0.0056421547,0.00085662724,0.0012628585,0.00079324405,0.0009874532,0.001358947],"category_scores_gemma":[0.0072735264,0.00037775838,0.0009583644,0.0026383058,0.00039240392,0.0032524557,0.0013961551,0.0013653506,0.0016435097],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010488379,0.0005157696,0.022701846,0.0012315926,0.00048474877,0.00043174357,0.0013327866,0.009502917,0.064817496,0.0031360704,0.009566467,0.88522977],"study_design_scores_gemma":[0.00022315129,0.0013511461,0.068051755,0.00042168348,0.0011695839,0.0029394573,0.0020939438,0.7847331,0.068452604,0.023597168,0.04659983,0.0003665743],"about_ca_topic_score_codex":0.002880789,"about_ca_topic_score_gemma":0.0063889776,"teacher_disagreement_score":0.0056421547,"about_ca_system_score_codex":0.00041184927,"about_ca_system_score_gemma":0.0007996346,"threshold_uncertainty_score":0.007954121},"labels":[],"label_agreement":null},{"id":"W2603829285","doi":"","title":"Best Practices in Reasearch Methods - Guidelines for Translating RCT Findings into Practice","year":2007,"lang":"en","type":"article","venue":"Canadian Journal of Nursing Research","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Randomized controlled trial; Medicine; Psychology; Internal medicine","score_opus":0.597513765333551,"score_gpt":0.6316468776523951,"score_spread":0.03413311231884408,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2603829285","genre_codex":"methods","genre_gemma":"methods","domain_codex":"methods","domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":"methods","prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004644752,0.040133856,0.735952,0.061830163,0.008918635,0.121529326,0.00387903,0.006390787,0.016721455],"genre_scores_gemma":[0.009087473,0.004273547,0.911781,0.0021593524,0.00027219814,0.070387155,0.00055935275,0.0004388999,0.0010410099],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.20743804,0.6224671,0.12169657,0.0062993052,0.03989175,0.002207295],"domain_scores_gemma":[0.13789271,0.65335023,0.024995124,0.05591935,0.122576535,0.0052661155],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.67191607,0.006957112,0.013031654,0.022179725,0.007945921,0.023361953,0.018613694,0.015917359,0.013753958],"category_scores_gemma":[0.826599,0.008810792,0.012953243,0.019618932,0.01006438,0.0144025665,0.014734268,0.01953786,0.008732121],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015579974,0.0015614327,0.0055794246,0.0693941,0.0037850582,0.00070193474,0.028579418,0.003360369,0.0008850772,0.030026427,0.11409742,0.74047136],"study_design_scores_gemma":[0.013889583,0.0020724004,0.011802608,0.37542188,0.011363155,0.0018960191,0.018777493,0.02117296,0.012183853,0.21847741,0.31111813,0.0018246153],"about_ca_topic_score_codex":0.020481134,"about_ca_topic_score_gemma":0.042983957,"teacher_disagreement_score":0.32808393,"about_ca_system_score_codex":0.019780414,"about_ca_system_score_gemma":0.06605272,"threshold_uncertainty_score":0.40458596},"labels":[],"label_agreement":null},{"id":"W2604228733","doi":"","title":"Event Nugget Detection Task: UMBC systems","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Computer science; Task (project management); Sentence; Measure (data warehouse); Convolution (computer science); Word (group theory); Artificial intelligence; Convolutional neural network; Artificial neural network; Real-time computing; Pattern recognition (psychology); Data mining; Mathematics; Engineering","score_opus":0.007871750298804013,"score_gpt":0.230625804219854,"score_spread":0.22275405392105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2604228733","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46907318,0.0053463876,0.20346278,0.004304354,0.0023761902,0.0027566734,0.07055549,0.20905733,0.033067606],"genre_scores_gemma":[0.6209551,0.0005581924,0.22542605,0.0012695849,0.00040904002,0.0011684502,0.13073413,0.0035421124,0.01593731],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953923,0.0010891306,0.00039032177,0.001188667,0.0014156444,0.0005239469],"domain_scores_gemma":[0.99188834,0.0028642856,0.00040641104,0.0017767369,0.002433062,0.0006311202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045854785,0.002565446,0.0016080429,0.0031171783,0.0021675692,0.0026639996,0.0026305479,0.0030861294,0.008193732],"category_scores_gemma":[0.019015599,0.00064211333,0.0007179935,0.0036191223,0.00039118758,0.004228416,0.003203898,0.0024488554,0.008925368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031311475,0.0010840204,0.01619934,0.0013322626,0.00033821148,0.0013809461,0.0014619762,0.01682874,0.040868007,0.0036473619,0.37316877,0.54055923],"study_design_scores_gemma":[0.00048325036,0.0007393345,0.033322092,0.00020801928,0.00020303606,0.0013403371,0.0015295923,0.6867938,0.08120308,0.01217537,0.1817025,0.0002996962],"about_ca_topic_score_codex":0.031560563,"about_ca_topic_score_gemma":0.032456294,"teacher_disagreement_score":0.031560563,"about_ca_system_score_codex":0.0022753878,"about_ca_system_score_gemma":0.0017464843,"threshold_uncertainty_score":0.06275374},"labels":[],"label_agreement":null},{"id":"W2604963202","doi":"10.1609/aaai.v31i1.10931","title":"Distributed Negative Sampling for Word Embeddings","year":2017,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Word2vec; Word (group theory); Vocabulary; Computer science; Sampling (signal processing); Artificial intelligence; Core (optical fiber); Scale (ratio); Natural language processing; Theoretical computer science; Mathematics; Linguistics; Telecommunications; Geography; Cartography","score_opus":0.16769677131130994,"score_gpt":0.35517976344463,"score_spread":0.18748299213332004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2604963202","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016368933,0.00041240163,0.9805839,0.0003186194,0.00011139416,0.000074507865,0.00016619824,0.00090071425,0.0010633877],"genre_scores_gemma":[0.5614216,0.00060318754,0.42430478,0.0006182569,0.000504443,0.0006421009,0.0031425268,0.0005030489,0.008260176],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99772006,0.0009577001,0.00009094398,0.00054130517,0.0005200721,0.00016976902],"domain_scores_gemma":[0.99456984,0.0031089094,0.00029914657,0.0010056013,0.00082700513,0.00018946979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025045315,0.0012913371,0.0016548984,0.0009369476,0.00094601646,0.0014554046,0.0020109909,0.0010081067,0.0032668111],"category_scores_gemma":[0.01644401,0.00055624446,0.00057376834,0.0012688707,0.0016032181,0.0038729857,0.0026451452,0.0018237229,0.0015468104],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011610335,0.00037187795,0.004405352,0.0004343222,0.00013803435,0.0002620871,0.0003697078,0.33780643,0.012682117,0.19483,0.029306687,0.41823232],"study_design_scores_gemma":[0.000033794277,0.000044792094,0.00015919616,0.000008783849,0.000007643792,0.000055836175,0.00003169339,0.92562306,0.0016118321,0.0706693,0.0017436664,0.000010387597],"about_ca_topic_score_codex":0.0023044136,"about_ca_topic_score_gemma":0.003974861,"teacher_disagreement_score":0.0032668111,"about_ca_system_score_codex":0.0010687842,"about_ca_system_score_gemma":0.0010313438,"threshold_uncertainty_score":0.013245344},"labels":[],"label_agreement":null},{"id":"W2605905982","doi":"10.1007/978-3-319-57351-9_29","title":"Domain Adaptation for Detecting Mild Cognitive Impairment","year":2017,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of the Fraser Valley; University of British Columbia","funders":"","keywords":"Computer science; Domain adaptation; Adaptation (eye); Cognitive impairment; Dementia; Domain (mathematical analysis); Cognition; Natural language processing; Artificial intelligence; Machine learning; Psychology; Medicine; Disease; Psychiatry","score_opus":0.0395186085395018,"score_gpt":0.2785107478688961,"score_spread":0.23899213932939434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2605905982","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4793575,0.009497645,0.48550957,0.0005217121,0.0005077628,0.00032687787,0.0036469288,0.0062774913,0.014354542],"genre_scores_gemma":[0.8583869,0.0024764827,0.12858401,0.00018116969,0.00025225565,0.0001969138,0.0042908145,0.00024099112,0.005390563],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99969923,0.00008597102,0.000018223589,0.00007999278,0.000067347886,0.000049268256],"domain_scores_gemma":[0.99929106,0.00034444296,0.00004715232,0.00009348285,0.00017184817,0.00005186617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008131709,0.0005905165,0.0005264826,0.0010796341,0.00018672575,0.0006727102,0.00044975957,0.00043896938,0.001901978],"category_scores_gemma":[0.0026796397,0.00010214087,0.0004749593,0.0006953328,0.00014344249,0.0004731547,0.00066381716,0.0006236548,0.001417331],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010623847,0.0002344656,0.021394271,0.00019663858,0.00016386392,0.0001872972,0.00012489005,0.009240168,0.0432319,0.000547295,0.010220455,0.9133964],"study_design_scores_gemma":[0.00014753998,0.0010751358,0.2606718,0.0002275484,0.00063290104,0.0037082613,0.0008009543,0.6230589,0.074791096,0.013289532,0.02141972,0.00017663528],"about_ca_topic_score_codex":0.0018317177,"about_ca_topic_score_gemma":0.0019803687,"teacher_disagreement_score":0.001901978,"about_ca_system_score_codex":0.00016649935,"about_ca_system_score_gemma":0.0003129525,"threshold_uncertainty_score":0.006362796},"labels":[],"label_agreement":null},{"id":"W2606337962","doi":"","title":"Generalized Probabilistic Topic and Syntax Models for Natural Language Processing","year":2012,"lang":"en","type":"dissertation","venue":"The Atrium (University of Guelph)","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Syntax; Computer science; Probabilistic logic; Natural language processing; Linguistics; Artificial intelligence; Programming language; Philosophy","score_opus":0.020636575539729776,"score_gpt":0.23836168055126097,"score_spread":0.2177251050115312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2606337962","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004021196,0.001565403,0.9891733,0.0011872692,0.00009327501,0.000081471444,0.00072723336,0.00069165445,0.0024592087],"genre_scores_gemma":[0.33443168,0.005848699,0.6353701,0.0012147499,0.0012504341,0.0016727187,0.0044783424,0.0009286958,0.014804662],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966563,0.0016407692,0.0001899071,0.00074974674,0.0005794508,0.00018389706],"domain_scores_gemma":[0.9916198,0.0063929725,0.0006419781,0.0006148402,0.00060436054,0.0001259054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041611334,0.0014951627,0.0015107028,0.0030011553,0.00096446584,0.0034715105,0.0029203116,0.002180326,0.004876371],"category_scores_gemma":[0.015672768,0.0010223726,0.0032626977,0.0037493373,0.0024391946,0.0077883843,0.0021170932,0.003529853,0.0019478728],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006500042,0.00006474775,0.0014287417,0.00035755892,0.00022695665,0.00023352366,0.0009004673,0.23514052,0.0009586919,0.6799292,0.007325405,0.07336927],"study_design_scores_gemma":[0.000013214839,0.000013952927,0.00031757655,0.000032097654,0.000028396807,0.00007827708,0.000047052254,0.47186285,0.00012410316,0.52088964,0.006562741,0.000030000589],"about_ca_topic_score_codex":0.009322238,"about_ca_topic_score_gemma":0.008239039,"teacher_disagreement_score":0.009322238,"about_ca_system_score_codex":0.0024905163,"about_ca_system_score_gemma":0.0018228638,"threshold_uncertainty_score":0.022006452},"labels":[],"label_agreement":null},{"id":"W2606584033","doi":"10.1007/978-3-319-57351-9_10","title":"Investigating Citation Linkage with Machine Learning","year":2017,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Sentence; Citation; Linkage (software); Similarity (geometry); Artificial intelligence; Natural language processing; Learning to rank; Task (project management); Rank (graph theory); Domain (mathematical analysis); Regression; Information retrieval; Machine learning; Statistics; Ranking (information retrieval); Mathematics; World Wide Web","score_opus":0.02924486356172764,"score_gpt":0.25107939023284415,"score_spread":0.2218345266711165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2606584033","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35857308,0.042337716,0.52370465,0.0058100703,0.0019253909,0.0002691646,0.004699848,0.0036741633,0.05900599],"genre_scores_gemma":[0.8266295,0.007807215,0.14302228,0.00037358835,0.002400287,0.00025261566,0.005462806,0.0007066845,0.013345049],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9936101,0.0031669156,0.00042997507,0.0010762224,0.0014186425,0.00029815928],"domain_scores_gemma":[0.9299385,0.057953775,0.004187487,0.004013263,0.003153199,0.00075373834],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.008537647,0.00085301965,0.0017651225,0.015702633,0.0026818577,0.0074700736,0.0025773249,0.0025566006,0.008944774],"category_scores_gemma":[0.08520725,0.00076080114,0.0014095778,0.034467757,0.0010419496,0.009430279,0.003109369,0.002260458,0.0036966742],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005702511,0.00052180374,0.08418671,0.0015786446,0.0012482968,0.00069748383,0.0016923266,0.0353003,0.004400404,0.10501064,0.03642106,0.72837216],"study_design_scores_gemma":[0.00008551523,0.0001932611,0.02537614,0.00051390723,0.0008366227,0.0011179405,0.0012582729,0.47200656,0.00820003,0.4394562,0.050825942,0.0001296681],"about_ca_topic_score_codex":0.0013669247,"about_ca_topic_score_gemma":0.0016060392,"teacher_disagreement_score":0.9842974,"about_ca_system_score_codex":0.0010992121,"about_ca_system_score_gemma":0.0013503883,"threshold_uncertainty_score":0.04515195},"labels":[],"label_agreement":null},{"id":"W2607475845","doi":"10.1007/978-3-319-57351-9_46","title":"Deep Multi-cultural Graph Representation Learning","year":2017,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Jaccard index; Similarity (geometry); Graph; Representation (politics); German; Knowledge representation and reasoning; Information retrieval; Pattern recognition (psychology); Theoretical computer science; Linguistics","score_opus":0.04364916223908568,"score_gpt":0.2978275706104297,"score_spread":0.25417840837134403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2607475845","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02486476,0.0049348306,0.9483993,0.0014019926,0.0003912443,0.00004694889,0.0009625284,0.003412606,0.015585735],"genre_scores_gemma":[0.5108972,0.006612881,0.42582723,0.0007948898,0.00039124727,0.00012564012,0.008010176,0.0009033415,0.046437345],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998053,0.00006070761,0.0000072920484,0.000066799206,0.00003328086,0.000026771853],"domain_scores_gemma":[0.9995939,0.0001615619,0.000025734636,0.00012641883,0.000059393424,0.00003301703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0003570804,0.00094022683,0.0006384145,0.0009828213,0.00038908218,0.0011925522,0.0013503773,0.0012632181,0.0054573617],"category_scores_gemma":[0.0015758248,0.000429043,0.0010792362,0.0019219677,0.00048086958,0.0024431997,0.0012724518,0.0022880407,0.0026740194],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000071198185,0.00012280798,0.0010060822,0.00028775434,0.00017927871,0.00010342534,0.00020888606,0.116824925,0.006458621,0.063914075,0.041984677,0.7688383],"study_design_scores_gemma":[0.000009681153,0.000031346783,0.00050205196,0.00006467395,0.00004833989,0.00011681126,0.00007220952,0.8817832,0.0027767448,0.10026906,0.014305568,0.000020373594],"about_ca_topic_score_codex":0.0052073617,"about_ca_topic_score_gemma":0.012619392,"teacher_disagreement_score":0.0054573617,"about_ca_system_score_codex":0.0008399912,"about_ca_system_score_gemma":0.00051654375,"threshold_uncertainty_score":0.018256724},"labels":[],"label_agreement":null},{"id":"W2608787653","doi":"10.18653/v1/p17-1152","title":"Enhanced LSTM for Natural Language Inference","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1196,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; York University; National Research Council Canada","funders":"Fundamental Research Funds for the Central Universities; Chinese Academy of Sciences","keywords":"Inference; Computer science; Artificial intelligence; Parsing; Language model; Machine learning; Natural language; Artificial neural network; Natural language processing","score_opus":0.03843647516708471,"score_gpt":0.3355466537017409,"score_spread":0.2971101785346562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2608787653","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018744977,0.0024020183,0.9557861,0.0011892583,0.00037021126,0.00007172611,0.0029072864,0.013484959,0.005043492],"genre_scores_gemma":[0.5102138,0.0020657184,0.4638233,0.00085227547,0.00048158615,0.0002700006,0.00974422,0.0009083831,0.01164064],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99938047,0.00015004493,0.000038079786,0.00023776115,0.00013094628,0.000062666724],"domain_scores_gemma":[0.9988174,0.0006645861,0.000070094175,0.00022307728,0.0001898597,0.000034824145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009987364,0.0012028058,0.00068586087,0.0010162527,0.00038237736,0.0011223161,0.0017061513,0.0013104822,0.0073193405],"category_scores_gemma":[0.004881216,0.0004920878,0.0009851548,0.0016060575,0.0004979695,0.004287284,0.0012451872,0.002996072,0.0036686114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029563505,0.00022370127,0.0013359761,0.0006265444,0.00026488552,0.00025629846,0.00022993449,0.25373158,0.025286922,0.03523801,0.034944873,0.6475657],"study_design_scores_gemma":[0.000009804585,0.000019953914,0.00019869496,0.00002159179,0.000029869678,0.000040134553,0.000013417315,0.9536819,0.0043685474,0.036984734,0.004620442,0.000010915463],"about_ca_topic_score_codex":0.0070216963,"about_ca_topic_score_gemma":0.013080575,"teacher_disagreement_score":0.0073193405,"about_ca_system_score_codex":0.0012330414,"about_ca_system_score_gemma":0.0014128732,"threshold_uncertainty_score":0.024485588},"labels":[],"label_agreement":null},{"id":"W2610350176","doi":"10.1109/tkde.2017.2699965","title":"Finding Related Forum Posts through Content Similarity over Intention-Based Segmentation","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Knowledge and Data Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"European Research Council; Université Paris-Saclay; European Cooperation in Science and Technology","keywords":"Segmentation; Similarity (geometry); Computer science; Set (abstract data type); Point (geometry); Information retrieval; Contrast (vision); Content (measure theory); Artificial intelligence; Image (mathematics); Mathematics","score_opus":0.0766944298027059,"score_gpt":0.3090613877724966,"score_spread":0.23236695796979068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2610350176","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3693159,0.0015287882,0.61994696,0.0005222565,0.0001218608,0.00079467247,0.00074172625,0.0020115895,0.005016286],"genre_scores_gemma":[0.74880683,0.0002998133,0.24610525,0.00010437786,0.00033748045,0.00035093,0.0014949508,0.00018952136,0.0023109156],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99762243,0.0006037924,0.00018634587,0.00072712934,0.000625586,0.00023467847],"domain_scores_gemma":[0.9888686,0.006355958,0.0018824242,0.0007892225,0.0015362119,0.0005674574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026743985,0.0011493608,0.0016756409,0.009365087,0.0014417128,0.0022693174,0.0014064172,0.0017829593,0.0019285604],"category_scores_gemma":[0.013481594,0.00041452984,0.0011529729,0.0052694385,0.0011777933,0.0043122163,0.0018415966,0.000996811,0.00089026283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021615038,0.0017923511,0.08143952,0.0012444154,0.00048921286,0.00083407335,0.0033566628,0.062615514,0.040160272,0.020548021,0.00765284,0.7777056],"study_design_scores_gemma":[0.00010475706,0.0011297709,0.034191303,0.000113000955,0.00027763937,0.0008167717,0.0014511484,0.8939607,0.017415049,0.04501634,0.0053866133,0.00013685953],"about_ca_topic_score_codex":0.0028437306,"about_ca_topic_score_gemma":0.004066708,"teacher_disagreement_score":0.009365087,"about_ca_system_score_codex":0.00095613673,"about_ca_system_score_gemma":0.0012449038,"threshold_uncertainty_score":0.014143765},"labels":[],"label_agreement":null},{"id":"W2611225023","doi":"10.1007/978-3-319-59041-7_13","title":"Supervised Methods to Support Online Scientific Data Triage","year":2017,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Université du Québec à Montréal","funders":"","keywords":"Triage; Computer science; Pipeline (software); Feature selection; Data science; Task (project management); Feature (linguistics); Supervised learning; Machine learning; Artificial intelligence; Engineering; Systems engineering; Medicine; Medical emergency","score_opus":0.10553585693560324,"score_gpt":0.36290682600492474,"score_spread":0.2573709690693215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2611225023","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00633979,0.00074554083,0.9828334,0.00028073872,0.00024689463,0.00012088501,0.0009536999,0.006740814,0.001738187],"genre_scores_gemma":[0.11029995,0.00084725145,0.8646376,0.0003597458,0.0006809629,0.0005037216,0.0074625956,0.0016205268,0.013587696],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984345,0.00046085642,0.00012340547,0.00037138155,0.0005053215,0.00010448982],"domain_scores_gemma":[0.99507195,0.0020848005,0.00034500397,0.0011222734,0.0012015287,0.00017445342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002135343,0.001053358,0.0010459492,0.0017574565,0.00078783225,0.0015337674,0.0023323826,0.0012426688,0.0048840516],"category_scores_gemma":[0.007191963,0.0007061737,0.0011094619,0.0019711207,0.00043895893,0.0025513335,0.0021021653,0.002476107,0.005622369],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028804495,0.0004605202,0.0024392344,0.0003181178,0.00022862895,0.000125136,0.00028484515,0.05508243,0.011621321,0.0091120545,0.07867494,0.8413647],"study_design_scores_gemma":[0.000027091646,0.000043022454,0.00077319064,0.00002936046,0.000040838015,0.00009312714,0.00006281576,0.9572371,0.0061199423,0.020622704,0.014928862,0.000021873342],"about_ca_topic_score_codex":0.0033303355,"about_ca_topic_score_gemma":0.0076751444,"teacher_disagreement_score":0.0048840516,"about_ca_system_score_codex":0.00051023066,"about_ca_system_score_gemma":0.0015220324,"threshold_uncertainty_score":0.016338825},"labels":[],"label_agreement":null},{"id":"W2615653823","doi":"10.71781/10122","title":"Narrative generation by associative network extraction from real-life temporal data","year":2016,"lang":"en","type":"dissertation","venue":"Open MIND","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fonds de recherche du Québec – Nature et technologies; Université de Montréal","keywords":"Narrative; Associative property; Extraction (chemistry); Computer science; Art; Literature; Chemistry; Mathematics","score_opus":0.10882618493417731,"score_gpt":0.37288356588304306,"score_spread":0.26405738094886577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2615653823","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08950739,0.0015127495,0.88664615,0.0006430479,0.00023875944,0.00040644492,0.0051896507,0.005612983,0.010242831],"genre_scores_gemma":[0.35193768,0.0013685346,0.6276545,0.000120943056,0.000093426985,0.0003792635,0.008233608,0.0003853676,0.009826686],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992354,0.0001638376,0.00007490609,0.00025543195,0.00022169655,0.00004874976],"domain_scores_gemma":[0.99782723,0.0011713271,0.00020529021,0.00030686441,0.00042718515,0.00006205701],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008851911,0.0008813894,0.0003923792,0.0026735882,0.0005877368,0.0015469312,0.00085552724,0.0006997395,0.0053779776],"category_scores_gemma":[0.006074314,0.0004182411,0.0007376155,0.002487534,0.00045771082,0.0024425257,0.0012394264,0.0006645138,0.0021162159],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064654066,0.00014627176,0.0064713443,0.0009232216,0.000109218156,0.0013552365,0.0021418012,0.02444277,0.031043908,0.015098973,0.010438923,0.90718174],"study_design_scores_gemma":[0.00007028825,0.00026745856,0.012865618,0.0004350772,0.00032791248,0.0018972247,0.0032465225,0.7218363,0.09603136,0.0505386,0.11232428,0.0001593052],"about_ca_topic_score_codex":0.004406512,"about_ca_topic_score_gemma":0.008823481,"teacher_disagreement_score":0.0053779776,"about_ca_system_score_codex":0.00074340205,"about_ca_system_score_gemma":0.0010284602,"threshold_uncertainty_score":0.017991126},"labels":[],"label_agreement":null},{"id":"W2618625790","doi":"","title":"Automatic Text and Speech Processing for the Detection of Dementia","year":2016,"lang":"en","type":"dissertation","venue":"TSpace (University of Toronto)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Alzheimer's Association","keywords":"Dementia; Speech recognition; Natural language processing; Computer science; Artificial intelligence; Psychology; Medicine; Pathology; Disease","score_opus":0.013310927963707827,"score_gpt":0.24205483310652975,"score_spread":0.22874390514282192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2618625790","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.563314,0.009819827,0.3621498,0.0022188413,0.00096793787,0.0013960028,0.015125659,0.02346618,0.02154176],"genre_scores_gemma":[0.7248531,0.0021311024,0.2515707,0.0005872079,0.00041667122,0.00055770436,0.011856222,0.00031455743,0.007712834],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989085,0.00029318783,0.00011808265,0.00027156094,0.00029906738,0.00010941981],"domain_scores_gemma":[0.9972512,0.0012690588,0.00032412284,0.00022745723,0.0008119184,0.000116238734],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016489092,0.0011072551,0.00063713023,0.003228454,0.0006100321,0.0014388128,0.00078951113,0.0011445469,0.004103909],"category_scores_gemma":[0.005333872,0.00028296624,0.00073178817,0.0010127418,0.00028523116,0.0010414014,0.0008369026,0.000739215,0.005960794],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001073913,0.00032940696,0.020557987,0.00074082793,0.00014844195,0.0008461456,0.00048620554,0.0034684094,0.103229605,0.0014875677,0.019327942,0.8483035],"study_design_scores_gemma":[0.00018157695,0.0017467065,0.25605014,0.0007286396,0.00046178562,0.0048211766,0.0016569335,0.40875733,0.25023302,0.01707807,0.05789934,0.00038533474],"about_ca_topic_score_codex":0.002171062,"about_ca_topic_score_gemma":0.0033925374,"teacher_disagreement_score":0.004103909,"about_ca_system_score_codex":0.00048121103,"about_ca_system_score_gemma":0.0008062922,"threshold_uncertainty_score":0.013728917},"labels":[],"label_agreement":null},{"id":"W2620728056","doi":"10.48550/arxiv.1708.03994","title":"Data Sets: Word Embeddings Learned from Tweets and General Data","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Word (group theory); Word embedding; Natural language processing; Embedding; Artificial intelligence; Representation (politics); Information retrieval; Sentiment analysis; Linguistics","score_opus":0.3216883880300847,"score_gpt":0.26253500969841354,"score_spread":0.05915337833167117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2620728056","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70251364,0.00268575,0.10568891,0.002328635,0.0010633379,0.0016074891,0.17180666,0.0051274057,0.0071782125],"genre_scores_gemma":[0.61204815,0.00085113547,0.11447025,0.0005163882,0.00027204273,0.001706808,0.26553375,0.0002567382,0.0043447046],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99840313,0.00048365904,0.00018925987,0.000494993,0.00029962618,0.0001294322],"domain_scores_gemma":[0.99455065,0.0023015912,0.00040967658,0.0016016779,0.0008897723,0.00024667577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017213017,0.0015091965,0.0006975606,0.00199517,0.00058259064,0.0009107565,0.0012301321,0.001946218,0.0030130106],"category_scores_gemma":[0.013556371,0.0004791942,0.0015856669,0.0028059965,0.000932887,0.0028406547,0.0016338084,0.0022766946,0.0021607524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049295877,0.0047239223,0.14929834,0.0037265846,0.0018074121,0.002587627,0.0013911915,0.14893784,0.021199508,0.008817229,0.16354442,0.48903632],"study_design_scores_gemma":[0.0007912722,0.0017228782,0.11374729,0.00030623176,0.00070045184,0.0029622535,0.0014494453,0.71257716,0.03714984,0.024730882,0.103429,0.00043328787],"about_ca_topic_score_codex":0.004169811,"about_ca_topic_score_gemma":0.0057570497,"teacher_disagreement_score":0.004169811,"about_ca_system_score_codex":0.0008049292,"about_ca_system_score_gemma":0.00062933407,"threshold_uncertainty_score":0.010079563},"labels":[],"label_agreement":null},{"id":"W2621376330","doi":"10.1162/tacl_a_00055","title":"Joint Modeling of Topics, Citations, and Topical Authority in Academic Corpora","year":2017,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science, ICT and Future Planning","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Citation; Search engine indexing; Information retrieval; Process (computing); Joint (building); Generative grammar; Data science; Artificial intelligence; World Wide Web","score_opus":0.07106014512581252,"score_gpt":0.318250592846475,"score_spread":0.24719044772066248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2621376330","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3560072,0.0077786786,0.6035477,0.0038060406,0.00045331192,0.0005694352,0.008518483,0.005805213,0.013513992],"genre_scores_gemma":[0.8158143,0.002586906,0.15588714,0.000314413,0.001016593,0.0009815642,0.015004647,0.0006796211,0.0077149044],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956542,0.0022157414,0.00034091348,0.0009961798,0.0005692671,0.00022366544],"domain_scores_gemma":[0.97457707,0.019741692,0.0015250313,0.0018261218,0.0018848637,0.0004451852],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.010390579,0.0011157229,0.0015515826,0.008389806,0.0016346027,0.004208154,0.0020734398,0.0021871626,0.0024341065],"category_scores_gemma":[0.045555748,0.000909063,0.0020142193,0.010194884,0.0016084472,0.0083046695,0.0022597138,0.002778608,0.0017647923],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008235031,0.0007243739,0.0446758,0.00097236596,0.00065774686,0.00067392236,0.0036620614,0.552632,0.0050151516,0.08699914,0.028907582,0.27425635],"study_design_scores_gemma":[0.000054784083,0.00003389429,0.0032311315,0.00004069344,0.0000602016,0.00008290669,0.00013184587,0.9552019,0.0006297481,0.03624069,0.004258548,0.000033607183],"about_ca_topic_score_codex":0.016475938,"about_ca_topic_score_gemma":0.0224506,"teacher_disagreement_score":0.99161017,"about_ca_system_score_codex":0.0023868072,"about_ca_system_score_gemma":0.0027237453,"threshold_uncertainty_score":0.05495131},"labels":[],"label_agreement":null},{"id":"W2622583534","doi":"","title":"Towards Abstractive Multi-Document Summarization Using Submodular Function-Based Framework, Sentence Compression and Merging","year":2016,"lang":"en","type":"article","venue":"International Joint Conference on Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Submodular set function; Automatic summarization; Computer science; Redundancy (engineering); Sentence; Artificial intelligence; Natural language processing; Scalability; Set (abstract data type); Multi-document summarization; Information retrieval; Mathematics; Database","score_opus":0.03583080388495406,"score_gpt":0.3049195049577958,"score_spread":0.2690887010728417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2622583534","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0061080754,0.00072685425,0.9904289,0.00014733017,0.000033627588,0.00008182824,0.00023519574,0.0019096876,0.00032845867],"genre_scores_gemma":[0.08466768,0.00066080317,0.9096042,0.00022557919,0.0002150083,0.00028425953,0.0023955528,0.0002963263,0.0016505858],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99824667,0.0006375524,0.00016094657,0.0004191165,0.0004422611,0.000093626986],"domain_scores_gemma":[0.99728143,0.0009212206,0.00040974014,0.0004609509,0.00081925176,0.000107287946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026801967,0.0020623412,0.0023717755,0.0035303747,0.00062777504,0.0019789585,0.001811266,0.0013015976,0.0012400233],"category_scores_gemma":[0.004376837,0.00044833997,0.0015158433,0.0032003468,0.00062403985,0.0026772516,0.0012607668,0.0016274469,0.0012005892],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024398853,0.00023608292,0.00090930064,0.0006925449,0.00028247075,0.00026690593,0.0004923813,0.13798851,0.045129973,0.013191303,0.011549634,0.78901696],"study_design_scores_gemma":[0.00003949524,0.00026417803,0.0006525848,0.000028342987,0.0001331239,0.00016423284,0.000103378945,0.9543264,0.020137934,0.016819634,0.00728412,0.000046574725],"about_ca_topic_score_codex":0.0020497835,"about_ca_topic_score_gemma":0.0025055537,"teacher_disagreement_score":0.0035303747,"about_ca_system_score_codex":0.0010362482,"about_ca_system_score_gemma":0.0012381155,"threshold_uncertainty_score":0.014174402},"labels":[],"label_agreement":null},{"id":"W2622780478","doi":"10.3758/s13428-017-0899-1","title":"AGSuite: Software to conduct feature analysis of artificial grammar learning performance","year":2017,"lang":"en","type":"article","venue":"Behavior Research Methods","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada; University of Manitoba","keywords":"Computer science; Artificial intelligence; Feature (linguistics); Natural language processing; Grammaticality; Grammar; Task (project management); Machine learning; String (physics); Levenshtein distance; Linguistics; Mathematics","score_opus":0.4500381688493872,"score_gpt":0.5853024771923007,"score_spread":0.13526430834291348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2622780478","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068980776,0.00006974911,0.56876135,0.00013729236,0.00012581509,0.0005848664,0.015561091,0.3408711,0.0049078967],"genre_scores_gemma":[0.34069014,0.00011973164,0.58831143,0.00013223618,0.00009701186,0.005061279,0.02213746,0.03300174,0.010448968],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993569,0.00019688309,0.00006159906,0.00018376671,0.00014898667,0.000051845942],"domain_scores_gemma":[0.99425113,0.004562219,0.0002638361,0.00039420673,0.00042298864,0.00010558844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018403946,0.0011351613,0.0006898814,0.002047289,0.00032607155,0.0009928585,0.0010617888,0.000584409,0.021078838],"category_scores_gemma":[0.0105356565,0.0004788317,0.0009183647,0.0008897942,0.00024014573,0.0011689891,0.0009526664,0.00090550864,0.005469412],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017698004,0.0010273923,0.046640486,0.0010655464,0.00091908063,0.000264365,0.0022538048,0.027063861,0.020611743,0.010078562,0.1740652,0.71424013],"study_design_scores_gemma":[0.0003791347,0.0007295677,0.04796976,0.00010470241,0.00031442672,0.00032441836,0.0003657531,0.83749485,0.03514326,0.019879935,0.05713966,0.00015446988],"about_ca_topic_score_codex":0.0026417898,"about_ca_topic_score_gemma":0.0035229465,"teacher_disagreement_score":0.021078838,"about_ca_system_score_codex":0.0004323713,"about_ca_system_score_gemma":0.0008362896,"threshold_uncertainty_score":0.07051569},"labels":[],"label_agreement":null},{"id":"W2624261825","doi":"10.1146/annurev-statistics-031017-100547","title":"Words, Words, Words: How the Digital Humanities Are Integrating Diverse Research Fields to Study People","year":2018,"lang":"en","type":"article","venue":"Annual Review of Statistics and Its Application","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Digital humanities; Scholarship; Mainstream; Field (mathematics); Big data; Interdisciplinarity; Data science; Digital scholarship; Political science; Sociology; Public relations; Social science; Library science; Computer science","score_opus":0.06264499503292599,"score_gpt":0.3498613314029844,"score_spread":0.28721633637005844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2624261825","genre_codex":"commentary","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038399454,0.17094351,0.12388467,0.48284915,0.012969023,0.00019449783,0.0019202454,0.0007102727,0.16812912],"genre_scores_gemma":[0.6831004,0.1010338,0.08047354,0.06676959,0.0143384095,0.00052944536,0.0011139954,0.0015222983,0.05111854],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98832357,0.007908103,0.00033460738,0.0012136054,0.0018844197,0.00033567328],"domain_scores_gemma":[0.97125274,0.021282269,0.0013309321,0.0029877867,0.0018698178,0.0012765522],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.010153001,0.0007732718,0.00077657157,0.007186633,0.005691792,0.021883773,0.001356084,0.0033978329,0.00947878],"category_scores_gemma":[0.032411844,0.0005119492,0.0005843388,0.009992179,0.037018005,0.050498437,0.0077979416,0.00565085,0.003042499],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007606357,0.000030954892,0.0038421622,0.00057583343,0.00006702415,0.0002213955,0.06431086,0.0002993753,0.00045342915,0.75199354,0.07149458,0.1066348],"study_design_scores_gemma":[0.00001139568,0.0000236281,0.001843297,0.00075285666,0.000023161127,0.0005881387,0.028274307,0.00043227643,0.00022084857,0.47923657,0.48853955,0.00005394597],"about_ca_topic_score_codex":0.005371618,"about_ca_topic_score_gemma":0.0071944515,"teacher_disagreement_score":0.99430823,"about_ca_system_score_codex":0.0031343456,"about_ca_system_score_gemma":0.0042268103,"threshold_uncertainty_score":0.053694844},"labels":[],"label_agreement":null},{"id":"W2626154462","doi":"10.18653/v1/k17-1028","title":"Making Neural QA as Simple as Possible but not Simpler","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":176,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Computer science; Simple (philosophy); Artificial intelligence; Task (project management); Artificial neural network; Context (archaeology); Question answering; Heuristic; Baseline (sea); Perspective (graphical); Machine learning; Engineering","score_opus":0.09895202575389195,"score_gpt":0.3529582687877168,"score_spread":0.25400624303382485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2626154462","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08181901,0.002041447,0.88338846,0.004600189,0.00035208848,0.00022660637,0.0010399685,0.013532576,0.012999687],"genre_scores_gemma":[0.56674546,0.0010840814,0.41836676,0.0015405507,0.0002911148,0.00027248345,0.0030683216,0.00061645283,0.008014816],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983575,0.00064016756,0.00009554132,0.00056582235,0.00024794653,0.00009309165],"domain_scores_gemma":[0.9937192,0.003005367,0.0002694041,0.0018542954,0.0009770448,0.00017469276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038519646,0.00065341656,0.00090353715,0.0006556576,0.00075630244,0.002797138,0.0018129649,0.0018301253,0.006695984],"category_scores_gemma":[0.017287437,0.00051988504,0.00065364514,0.0007154929,0.0010827394,0.008910659,0.0021969697,0.0032803735,0.004318256],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006456325,0.00038606487,0.0044441987,0.0012990261,0.000287463,0.00018892005,0.001054316,0.09095102,0.05997955,0.07843065,0.030712891,0.7316203],"study_design_scores_gemma":[0.00006818474,0.0002526959,0.0021873044,0.00012668876,0.000122492,0.00026803825,0.0002946314,0.76237345,0.02880186,0.17714487,0.028299112,0.000060751216],"about_ca_topic_score_codex":0.0031725862,"about_ca_topic_score_gemma":0.0051338556,"teacher_disagreement_score":0.006695984,"about_ca_system_score_codex":0.00082204986,"about_ca_system_score_gemma":0.0010142786,"threshold_uncertainty_score":0.02240032},"labels":[],"label_agreement":null},{"id":"W2727955266","doi":"10.1145/3079628.3079700","title":"Encoding User as More Than the Sum of Their Parts","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; TUTOR; Matching (statistics); Encoding (memory); Context (archaeology); Word (group theory); Recall; Artificial intelligence; Artificial neural network; Natural language processing; Linguistics; Programming language","score_opus":0.05777723805927694,"score_gpt":0.2870989742108666,"score_spread":0.22932173615158968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2727955266","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16244748,0.0022804867,0.8197935,0.0011923072,0.00040705624,0.0002654419,0.0054429257,0.0026475384,0.005523238],"genre_scores_gemma":[0.7601203,0.0013007453,0.22119641,0.0003327974,0.00024900047,0.00035198516,0.006273453,0.00025771168,0.009917608],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881554,0.0003606022,0.00011639321,0.00041313723,0.00020649821,0.00008787713],"domain_scores_gemma":[0.997322,0.0013406054,0.00020967046,0.0007032514,0.00033737512,0.000086967884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010812923,0.0011337607,0.0006621255,0.0018158755,0.0002701673,0.0016637983,0.00076136866,0.0011134996,0.0052638925],"category_scores_gemma":[0.0057819267,0.0004211662,0.00091765175,0.0023938287,0.0005374987,0.006094612,0.0012756694,0.0010852722,0.002061966],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017109177,0.00077080604,0.041573662,0.0008142417,0.00053595164,0.0004974026,0.001439263,0.04672171,0.024969866,0.045460474,0.019943513,0.81556225],"study_design_scores_gemma":[0.00006099848,0.00040181214,0.012585943,0.00011856268,0.00039147426,0.0007745028,0.00068931055,0.8288875,0.01157102,0.11985902,0.024537627,0.00012222724],"about_ca_topic_score_codex":0.0032825836,"about_ca_topic_score_gemma":0.0052888067,"teacher_disagreement_score":0.0052638925,"about_ca_system_score_codex":0.00053421734,"about_ca_system_score_gemma":0.0004974384,"threshold_uncertainty_score":0.017609477},"labels":[],"label_agreement":null},{"id":"W2729052508","doi":"10.1111/coin.12120","title":"From French Wikipedia to Erudit: A test case for cross‐domain open information extraction","year":2017,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Classifier (UML); Pipeline (software); Information extraction; Entity linking; Information retrieval; Open domain; Domain (mathematical analysis); Natural language processing; Task (project management); Artificial intelligence; Named-entity recognition; Precision and recall; Question answering; Knowledge base; Mathematics; Programming language","score_opus":0.07660808466055637,"score_gpt":0.39096216997523464,"score_spread":0.3143540853146783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2729052508","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8236721,0.0038803532,0.093480766,0.0053081354,0.00084295264,0.0008970392,0.017040193,0.026403053,0.028475441],"genre_scores_gemma":[0.83717674,0.000645937,0.11929394,0.0013570361,0.000215967,0.00041458604,0.032635067,0.0015166851,0.006744047],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99345917,0.002866832,0.0005824695,0.0011380535,0.0015250263,0.00042850192],"domain_scores_gemma":[0.97601753,0.015621276,0.0006903685,0.003413766,0.0037089947,0.0005480279],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051793465,0.0010597359,0.0008039875,0.0046474542,0.0020269773,0.0026170996,0.0011304473,0.002621264,0.0031048479],"category_scores_gemma":[0.02441605,0.00030027493,0.0007738206,0.0029881487,0.0010991556,0.0038391652,0.0039290385,0.00131667,0.002498297],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002208109,0.0022150239,0.06954047,0.0034925765,0.0006190834,0.02310375,0.013733193,0.018207962,0.03309835,0.0143196145,0.13090667,0.6885552],"study_design_scores_gemma":[0.0006717235,0.0013567985,0.06237429,0.0009468114,0.0004530386,0.016804075,0.012849322,0.21487266,0.15662874,0.02302644,0.50961024,0.00040581237],"about_ca_topic_score_codex":0.01360348,"about_ca_topic_score_gemma":0.013190681,"teacher_disagreement_score":0.01360348,"about_ca_system_score_codex":0.0010713781,"about_ca_system_score_gemma":0.0011997666,"threshold_uncertainty_score":0.027391315},"labels":[],"label_agreement":null},{"id":"W2734812719","doi":"10.18653/v1/w17-4308","title":"Piecewise Latent Variables for Neural Variational Text Processing","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Nuance Foundation; Canadian Institute for Advanced Research","keywords":"Latent variable; Computer science; Piecewise; Artificial intelligence; Prior probability; Latent variable model; Probabilistic latent semantic analysis; Inference; Exponential family; Gaussian; Natural language; Machine learning; Algorithm; Mathematics; Bayesian probability","score_opus":0.06619283798847073,"score_gpt":0.29702485725858024,"score_spread":0.23083201927010952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2734812719","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002230361,0.0004678376,0.99505854,0.00042045835,0.000041752144,0.000019171024,0.00014314103,0.00024042172,0.001378306],"genre_scores_gemma":[0.35273743,0.0025460594,0.6199787,0.000723891,0.00057589804,0.000625958,0.0015785056,0.001120426,0.02011315],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992873,0.00032366405,0.00003132472,0.0001595617,0.00014404947,0.00005414894],"domain_scores_gemma":[0.99762565,0.0017550843,0.00016369973,0.00020374474,0.00016341057,0.00008847563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015988651,0.0009075877,0.00076531013,0.0009470463,0.0005676386,0.001379894,0.0019112055,0.0015127906,0.006894854],"category_scores_gemma":[0.008203669,0.0007482978,0.0010351967,0.0013828125,0.0014693765,0.0031806508,0.0019960715,0.0036606707,0.0013936231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005391711,0.000029849893,0.00036691964,0.00015037587,0.00005459337,0.00006534232,0.00015385474,0.32140446,0.0016945528,0.63196576,0.004729071,0.03933134],"study_design_scores_gemma":[0.000004267671,0.0000045717097,0.000048128477,0.0000089421,0.000003850757,0.000010108347,0.0000061453347,0.78027475,0.00021563457,0.21765469,0.001762628,0.000006356356],"about_ca_topic_score_codex":0.005281141,"about_ca_topic_score_gemma":0.005690219,"teacher_disagreement_score":0.006894854,"about_ca_system_score_codex":0.0019072817,"about_ca_system_score_gemma":0.0011356274,"threshold_uncertainty_score":0.023065567},"labels":[],"label_agreement":null},{"id":"W2735682218","doi":"10.5539/ijel.v7n4p33","title":"Addressing the Problem of Coherence in Automatic Text Summarization: A Latent Semantic Analysis Approach","year":2017,"lang":"en","type":"article","venue":"International Journal of English Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Latent semantic analysis; Coherence (philosophical gambling strategy); Information retrieval; Natural language processing; Multi-document summarization; Semantics (computer science); Artificial intelligence; Statistics; Mathematics","score_opus":0.06497002097453258,"score_gpt":0.31786455789884255,"score_spread":0.25289453692430997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2735682218","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026478987,0.001100242,0.9704645,0.0004912479,0.00003242759,0.000111934794,0.00017729928,0.00046582622,0.00067746267],"genre_scores_gemma":[0.42740428,0.0010668205,0.56857634,0.00011061521,0.00023455324,0.00035704268,0.0011288078,0.00010088029,0.0010206656],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973356,0.0013852675,0.00027652,0.0004286456,0.00046772108,0.00010619649],"domain_scores_gemma":[0.99148107,0.0052049635,0.0013522495,0.00048434807,0.0013617848,0.00011558436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003920432,0.00071817933,0.0009516648,0.0052412376,0.00087901135,0.00208612,0.0008014877,0.00074659294,0.0008135255],"category_scores_gemma":[0.012208684,0.0003317486,0.0011249385,0.0030182193,0.0008800661,0.0036245016,0.0013745747,0.00099568,0.00033114792],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038484327,0.00030143213,0.0071863574,0.0014943106,0.00050369103,0.00030789326,0.004068304,0.037249207,0.030324472,0.05520357,0.004239963,0.858736],"study_design_scores_gemma":[0.0001131243,0.00061697746,0.012416675,0.00028983192,0.0006053355,0.00034216302,0.0021837673,0.8081239,0.020162832,0.14054297,0.014423452,0.0001789836],"about_ca_topic_score_codex":0.0012053249,"about_ca_topic_score_gemma":0.0015144295,"teacher_disagreement_score":0.0052412376,"about_ca_system_score_codex":0.0007008771,"about_ca_system_score_gemma":0.0014075722,"threshold_uncertainty_score":0.020733476},"labels":[],"label_agreement":null},{"id":"W2736532806","doi":"10.1007/s00799-017-0225-7","title":"The context of multiple in-text references and their signification","year":2017,"lang":"en","type":"article","venue":"International Journal on Digital Libraries","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Rhetorical question; Context (archaeology); Section (typography); Noun phrase; Point (geometry); Natural language processing; Information retrieval; Linguistics; Artificial intelligence; Noun; Mathematics; History","score_opus":0.03906275880157609,"score_gpt":0.2561606087268078,"score_spread":0.2170978499252317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2736532806","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39304397,0.01177718,0.17932314,0.012780841,0.0029546237,0.00023560603,0.0020620334,0.0016793528,0.3961433],"genre_scores_gemma":[0.9520754,0.0015726758,0.026348392,0.00019468732,0.0010504491,0.00008587955,0.00066114037,0.000634341,0.017377099],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952743,0.002021058,0.00035842787,0.0007447213,0.0011879987,0.00041341106],"domain_scores_gemma":[0.95684135,0.028273694,0.0036739812,0.0033447957,0.0065504164,0.0013157489],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.003142216,0.0006652739,0.00061821355,0.011724832,0.0051904907,0.011977537,0.0013472063,0.0018123232,0.012323476],"category_scores_gemma":[0.04294387,0.0006494518,0.00041321092,0.012372662,0.0033826884,0.0168249,0.0050986405,0.0022232821,0.0025488497],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076930475,0.00011053573,0.009557572,0.0010023382,0.00006842093,0.0023122157,0.09589838,0.0007628817,0.014636218,0.72340566,0.010384867,0.14109159],"study_design_scores_gemma":[0.00010749087,0.00030305274,0.021664836,0.0014870474,0.00045228878,0.0021051983,0.06477666,0.011078089,0.02103536,0.4404708,0.4363337,0.00018542362],"about_ca_topic_score_codex":0.0020963207,"about_ca_topic_score_gemma":0.0020635563,"teacher_disagreement_score":0.99685776,"about_ca_system_score_codex":0.0022765452,"about_ca_system_score_gemma":0.0015959075,"threshold_uncertainty_score":0.04122609},"labels":[],"label_agreement":null},{"id":"W2737301218","doi":"","title":"Query Expansion Using Pseudo Relevance Feedback on Wikipedia","year":2016,"lang":"en","type":"article","venue":"Web Search and Data Mining","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Athabasca University","funders":"","keywords":"Relevance feedback; Computer science; Query expansion; Relevance (law); Information retrieval; World Wide Web; Artificial intelligence; Political science","score_opus":0.11441356449461505,"score_gpt":0.32874768043468955,"score_spread":0.2143341159400745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2737301218","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26703876,0.0059506143,0.6995599,0.0016423297,0.00093827874,0.0007968302,0.003033212,0.011314398,0.009725607],"genre_scores_gemma":[0.77061146,0.00087241374,0.21537006,0.00030748494,0.00053635484,0.00032184622,0.0040093334,0.00044061194,0.007530477],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99750286,0.0010334702,0.00017304771,0.00035925524,0.0007623814,0.00016910533],"domain_scores_gemma":[0.993141,0.004086787,0.00017265287,0.00052496226,0.0018965201,0.00017811646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018966356,0.0009965023,0.0014781781,0.003595929,0.0008029473,0.0011059669,0.0010686218,0.0009844506,0.003491226],"category_scores_gemma":[0.012441319,0.0004229163,0.00083056797,0.0022216241,0.00036611946,0.0026016207,0.0010112942,0.00079889974,0.0015740882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026571944,0.0011251494,0.0054841656,0.0012891307,0.00027870553,0.0007241332,0.0005645712,0.071171366,0.06716685,0.007850572,0.05073577,0.7909523],"study_design_scores_gemma":[0.0001047905,0.0002595614,0.0015836579,0.000034001718,0.00013419717,0.0002901358,0.00010626626,0.9738931,0.013298596,0.0059940377,0.004256326,0.00004532583],"about_ca_topic_score_codex":0.0060420525,"about_ca_topic_score_gemma":0.011274862,"teacher_disagreement_score":0.0060420525,"about_ca_system_score_codex":0.00063053396,"about_ca_system_score_gemma":0.0014739106,"threshold_uncertainty_score":0.012013793},"labels":[],"label_agreement":null},{"id":"W2739730962","doi":"10.18653/v1/s17-1026","title":"Generating Pattern-Based Entailment Graphs for Relation Extraction","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Logical consequence; Relation (database); Computer science; Textual entailment; Relationship extraction; Structuring; Artificial intelligence; Natural language processing; Task (project management); Meaning (existential); Exploit; Semantic relation; Data mining","score_opus":0.05132058273999077,"score_gpt":0.3067872736571209,"score_spread":0.25546669091713015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2739730962","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017311133,0.0003388324,0.9706325,0.00030373843,0.000039930557,0.0003645631,0.002226346,0.006979941,0.0018029206],"genre_scores_gemma":[0.08501075,0.00035398806,0.902473,0.0001223686,0.00003507056,0.00028870368,0.010020033,0.0007729034,0.0009231314],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975262,0.0008582028,0.00026146538,0.00063965283,0.0006025955,0.00011191297],"domain_scores_gemma":[0.99258304,0.0041230675,0.0005053396,0.0017956828,0.000883157,0.0001096585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019035301,0.001638201,0.00087215647,0.0052513336,0.0011140646,0.001491209,0.0019085639,0.0014997848,0.004159124],"category_scores_gemma":[0.016113726,0.0007016199,0.0020860168,0.004683656,0.0007563662,0.00508305,0.0020654367,0.0015390217,0.0029656715],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003575477,0.0005013265,0.0049584964,0.0018008604,0.00035760528,0.0011695552,0.0012823252,0.05033906,0.048817046,0.07260331,0.025322635,0.79249024],"study_design_scores_gemma":[0.00009939103,0.00014884405,0.0023261732,0.00016759009,0.00023201988,0.0011733659,0.00040936895,0.67500854,0.06784506,0.21460278,0.037905,0.00008176603],"about_ca_topic_score_codex":0.0027124279,"about_ca_topic_score_gemma":0.0061175833,"teacher_disagreement_score":0.0052513336,"about_ca_system_score_codex":0.0010267358,"about_ca_system_score_gemma":0.0016107552,"threshold_uncertainty_score":0.013913631},"labels":[],"label_agreement":null},{"id":"W2740006839","doi":"10.18653/v1/p17-1114","title":"A Local Detection Approach for Named Entity Recognition and Mention Detection","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":143,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Named-entity recognition; Sequence labeling; Task (project management); ENCODE; Sentence; Artificial intelligence; Fragment (logic); Encoding (memory); Forgetting; Sequence (biology); Representation (politics); Natural language processing; Pattern recognition (psychology); Algorithm","score_opus":0.06030778319924279,"score_gpt":0.2589302997170373,"score_spread":0.1986225165177945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2740006839","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065354696,0.00036455068,0.98926663,0.00007922521,0.000055434684,0.00004026864,0.00013573607,0.0028560776,0.0006665997],"genre_scores_gemma":[0.17863348,0.000592504,0.80875534,0.00026011033,0.00025318918,0.00016982437,0.0019878088,0.00053723995,0.00881049],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99846077,0.00032495946,0.00008593554,0.0007006451,0.00032002974,0.00010767102],"domain_scores_gemma":[0.9963201,0.0014112183,0.0003468392,0.0010503524,0.0007305639,0.00014089691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015317211,0.0011829282,0.0011076909,0.0027814223,0.00071165443,0.0011048918,0.0022814027,0.0013671215,0.003276184],"category_scores_gemma":[0.0036368887,0.00050756324,0.0013267385,0.0023355903,0.00078231975,0.004078698,0.0015037869,0.0018615607,0.002972574],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026665226,0.00024027869,0.0036372512,0.0003500154,0.00014939376,0.0002446876,0.00029528895,0.027059592,0.06143087,0.009480952,0.006645348,0.89019966],"study_design_scores_gemma":[0.000027543852,0.00022105791,0.0034063421,0.00005255535,0.00016468806,0.0007708523,0.0001523871,0.8748913,0.083000734,0.017697936,0.019513514,0.00010112683],"about_ca_topic_score_codex":0.0021308095,"about_ca_topic_score_gemma":0.005010144,"teacher_disagreement_score":0.003276184,"about_ca_system_score_codex":0.00069313863,"about_ca_system_score_gemma":0.0008327247,"threshold_uncertainty_score":0.010959864},"labels":[],"label_agreement":null},{"id":"W2741333084","doi":"10.18653/v1/p17-1103","title":"Towards an Automatic Turing Test: Learning to Evaluate Dialogue Responses","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Canadian Institute for Advanced Research; McGill University","funders":"","keywords":"Turing test; Computer science; Utterance; Artificial intelligence; Quality (philosophy); Machine learning; Word (group theory); Test (biology); Natural language processing; Linguistics","score_opus":0.0891649191737554,"score_gpt":0.35126035107422837,"score_spread":0.26209543190047296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2741333084","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27277252,0.0044444865,0.6387936,0.00445531,0.0029328323,0.002111137,0.0066100024,0.039727133,0.028152969],"genre_scores_gemma":[0.7342633,0.0004567349,0.23836306,0.0015381087,0.000560769,0.0022932074,0.0140740555,0.00174471,0.006705988],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.948053,0.039670505,0.0025646493,0.0055023073,0.003114544,0.0010949512],"domain_scores_gemma":[0.873756,0.08910325,0.0033456062,0.015522849,0.014146417,0.0041258154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029107396,0.0033986596,0.0019924215,0.0031875735,0.0014450833,0.006037198,0.004705087,0.0054283976,0.007817908],"category_scores_gemma":[0.16739574,0.0010472484,0.001436192,0.0013092377,0.0021546653,0.008713468,0.009239824,0.005807073,0.014667699],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004000738,0.0021252744,0.02982928,0.0017940624,0.00097597693,0.0002824504,0.0030745626,0.026418602,0.014177552,0.011288866,0.10895319,0.7970795],"study_design_scores_gemma":[0.00092758366,0.0022570118,0.013996882,0.00055613497,0.0003124893,0.00039928447,0.0027385466,0.826465,0.027641855,0.09389965,0.030519,0.0002866508],"about_ca_topic_score_codex":0.0015504417,"about_ca_topic_score_gemma":0.001926853,"teacher_disagreement_score":0.029107396,"about_ca_system_score_codex":0.0013761347,"about_ca_system_score_gemma":0.0017576589,"threshold_uncertainty_score":0.15393645},"labels":[],"label_agreement":null},{"id":"W2741544939","doi":"10.1145/3077136.3080667","title":"Finally, a Downloadable Test Collection of Tweets","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Microblogging; World Wide Web; The Internet; Scalability; Download; Social media; Data collection; Information retrieval; Data science; Database","score_opus":0.0305441985586857,"score_gpt":0.260896675642129,"score_spread":0.23035247708344328,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2741544939","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30721527,0.0014037446,0.0268721,0.0028819619,0.0026285518,0.0039316495,0.5883589,0.03311367,0.033594172],"genre_scores_gemma":[0.22298019,0.00039321676,0.04679303,0.0007426832,0.0005794963,0.0036975902,0.70654637,0.003377535,0.01488998],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934088,0.0014544727,0.00083550985,0.0011224779,0.0025387872,0.00064008305],"domain_scores_gemma":[0.97495663,0.0051136194,0.0012383135,0.0094749285,0.007942082,0.0012743702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035582108,0.0014671828,0.0011269335,0.004478552,0.002700376,0.0024713026,0.0018004461,0.0016264663,0.00885404],"category_scores_gemma":[0.021179019,0.000623225,0.0012348837,0.0061231493,0.0013108947,0.0035833558,0.002731928,0.0027068378,0.011112248],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027673175,0.0022377018,0.039959144,0.0027174365,0.00047780995,0.0016381519,0.0025744026,0.0051191645,0.024243848,0.0039538187,0.790215,0.124096274],"study_design_scores_gemma":[0.00058675907,0.0008978808,0.12315508,0.0003731747,0.000292973,0.0019692841,0.0031482126,0.031967632,0.06284064,0.004256713,0.7701888,0.00032282123],"about_ca_topic_score_codex":0.00928417,"about_ca_topic_score_gemma":0.019143712,"teacher_disagreement_score":0.00928417,"about_ca_system_score_codex":0.0013071962,"about_ca_system_score_gemma":0.0018160345,"threshold_uncertainty_score":0.029619694},"labels":[],"label_agreement":null},{"id":"W2741782787","doi":"10.29007/fm8f","title":"Overview of COLIEE 2017","year":2018,"lang":"en","type":"article","venue":"EPiC series in computing","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Core Research for Evolutional Science and Technology; Ministry of Education, Culture, Sports, Science and Technology; Alberta Machine Intelligence Institute","keywords":"Task (project management); Computer science; Logical consequence; Textual entailment; Information retrieval; Information extraction; Group (periodic table); Task group; Question answering; Natural language processing; Artificial intelligence","score_opus":0.08135760093089604,"score_gpt":0.33177475092000414,"score_spread":0.2504171499891081,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2741782787","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060666878,0.067638524,0.13376044,0.025367822,0.030360641,0.012125682,0.26259375,0.20778093,0.19970527],"genre_scores_gemma":[0.038959295,0.005305135,0.104405545,0.005562577,0.0037980753,0.0044135656,0.7802992,0.013879731,0.043376982],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96604276,0.010033994,0.002192189,0.0045986488,0.013428652,0.003703776],"domain_scores_gemma":[0.9643918,0.0065141246,0.00088016194,0.0067731626,0.015773596,0.005667115],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028978122,0.005845866,0.0044808784,0.017769033,0.005292596,0.013668068,0.010414809,0.0061076414,0.048843566],"category_scores_gemma":[0.045907527,0.0020652893,0.0038999263,0.012884654,0.0020468272,0.013192668,0.013376384,0.007318138,0.06528703],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089794694,0.0010796405,0.0015148458,0.0017933722,0.0002961533,0.0001617215,0.00020698072,0.0024086682,0.0033859706,0.0025570875,0.8830244,0.1026733],"study_design_scores_gemma":[0.00080589013,0.00068592536,0.0041623083,0.0004788668,0.00024244147,0.00050107646,0.00042570158,0.019262653,0.008314131,0.006548439,0.9583176,0.0002549488],"about_ca_topic_score_codex":0.03258332,"about_ca_topic_score_gemma":0.054525435,"teacher_disagreement_score":0.048843566,"about_ca_system_score_codex":0.0061165052,"about_ca_system_score_gemma":0.0131021505,"threshold_uncertainty_score":0.16339797},"labels":[],"label_agreement":null},{"id":"W2744909235","doi":"10.1016/j.neunet.2014.09.005","title":"Challenges in representation learning: A report on three machine learning contests","year":2014,"lang":"en","type":"article","venue":"Neural Networks","topic":"Topic Modeling","field":"Computer Science","cited_by":676,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto; Université de Montréal","funders":"","keywords":"Computer science; Representation (politics); Artificial intelligence; Learning to learn; Feature learning; Machine learning; Mathematics education; Psychology","score_opus":0.07473609516329448,"score_gpt":0.2917676883922016,"score_spread":0.21703159322890708,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2744909235","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16335835,0.16135286,0.22159903,0.34685698,0.02382824,0.001006739,0.008055005,0.0028768245,0.071066014],"genre_scores_gemma":[0.58762616,0.070321806,0.1869653,0.02348068,0.026215555,0.0012662683,0.034667525,0.0025696722,0.066887036],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9890695,0.0032017853,0.0007537503,0.001320839,0.0043216147,0.0013324833],"domain_scores_gemma":[0.9398355,0.03492837,0.0014372886,0.004708101,0.014150345,0.0049404465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03071552,0.0014790798,0.0024872534,0.0032308486,0.003719024,0.009159641,0.0032353848,0.004486783,0.00670192],"category_scores_gemma":[0.060791425,0.0005474741,0.0018807827,0.0056941,0.0025373865,0.011458158,0.008992145,0.0065659755,0.0029397851],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005149687,0.0005983409,0.0040131193,0.0012486095,0.00017857249,0.00024578787,0.0015706191,0.0056291833,0.0027025389,0.054329164,0.43507436,0.49389482],"study_design_scores_gemma":[0.00018140337,0.000369,0.009283717,0.00057033525,0.00014421034,0.00058838405,0.0046059024,0.038014762,0.006953949,0.14850448,0.7906026,0.00018124809],"about_ca_topic_score_codex":0.009163261,"about_ca_topic_score_gemma":0.009796759,"teacher_disagreement_score":0.03071552,"about_ca_system_score_codex":0.0036202162,"about_ca_system_score_gemma":0.0057404297,"threshold_uncertainty_score":0.1624412},"labels":[],"label_agreement":null},{"id":"W2751124354","doi":"10.48550/arxiv.1709.02349","title":"A Deep Reinforcement Learning Chatbot","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":200,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Chatbot; Reinforcement learning; Computer science; Artificial intelligence; Artificial neural network; Deep learning; Machine learning; Sequence (biology); Ensemble learning; Natural language processing","score_opus":0.09029315580694751,"score_gpt":0.19739711345255698,"score_spread":0.10710395764560947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2751124354","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11962142,0.0012717181,0.82646805,0.0016192703,0.00082545687,0.00064208347,0.0009174827,0.031312533,0.01732202],"genre_scores_gemma":[0.75715244,0.00017731666,0.21817331,0.0009259963,0.00010933459,0.0005670844,0.0010897177,0.00043769972,0.021367092],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994338,0.00020253196,0.000022976474,0.0001522885,0.00011403804,0.00007435615],"domain_scores_gemma":[0.99886477,0.00056290854,0.000063403284,0.00013248235,0.00017260182,0.00020386167],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011233863,0.00082741934,0.0007279793,0.00036267546,0.0006424219,0.0005947497,0.0021370722,0.001377978,0.007961829],"category_scores_gemma":[0.003768764,0.00034472844,0.00039418877,0.00023947237,0.00064199127,0.001349897,0.0016647079,0.0016450947,0.0020607165],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027974257,0.00289724,0.007236251,0.0011028484,0.00036851174,0.0014756138,0.0010230757,0.36162096,0.057748117,0.040610373,0.06444249,0.458677],"study_design_scores_gemma":[0.00013966335,0.00028381866,0.0003742671,0.000023870714,0.000024457891,0.00013635987,0.00003456541,0.97582793,0.005902863,0.008834259,0.008387657,0.000030372137],"about_ca_topic_score_codex":0.0033519971,"about_ca_topic_score_gemma":0.0034731424,"teacher_disagreement_score":0.007961829,"about_ca_system_score_codex":0.0006977979,"about_ca_system_score_gemma":0.0009358458,"threshold_uncertainty_score":0.026634991},"labels":[],"label_agreement":null},{"id":"W2751220357","doi":"","title":"Recurrent Normalization Propagation","year":2017,"lang":"en","type":"article","venue":"International Conference on Learning Representations","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Normalization (sociology); Initialization; Computer science; Computation; Parametrization (atmospheric modeling); Artificial intelligence; Generative grammar; Machine learning; Algorithm","score_opus":0.08877434203423155,"score_gpt":0.3726418993060283,"score_spread":0.28386755727179674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2751220357","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065724845,0.00026715014,0.9791981,0.0003009608,0.00028593017,0.000077931516,0.00037319792,0.008247156,0.0046770955],"genre_scores_gemma":[0.40304533,0.00079065294,0.553899,0.0008890788,0.00040981505,0.00054111175,0.0023165143,0.0029943676,0.03511405],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992586,0.00010845182,0.000040800118,0.00026038755,0.0002282296,0.000103434635],"domain_scores_gemma":[0.9987509,0.00035006157,0.000098251476,0.00032960737,0.00042223255,0.000049035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011919768,0.0019842254,0.00110913,0.0008259313,0.00065931893,0.00177165,0.0028849274,0.001371375,0.012897773],"category_scores_gemma":[0.006130673,0.0007913237,0.0011460821,0.001098674,0.0008452949,0.0038128798,0.0017292564,0.0024438063,0.0063873674],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029295575,0.00019279492,0.001427271,0.00024517564,0.00018578857,0.00027601127,0.00028358144,0.23618203,0.038707748,0.053496186,0.032332245,0.6363781],"study_design_scores_gemma":[0.000021507942,0.000046686466,0.00026261053,0.000024383036,0.00004781512,0.00010111307,0.000027050348,0.9441276,0.019651122,0.026023267,0.009636305,0.000030504068],"about_ca_topic_score_codex":0.007774104,"about_ca_topic_score_gemma":0.012401976,"teacher_disagreement_score":0.012897773,"about_ca_system_score_codex":0.0010008175,"about_ca_system_score_gemma":0.001606196,"threshold_uncertainty_score":0.043147326},"labels":[],"label_agreement":null},{"id":"W2752083635","doi":"10.1145/3103010.3121039","title":"High-performance Computational Framework for Phrase Relatedness","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; n-gram; Phrase; Semantic similarity; Hash function; Encoding (memory); Similarity (geometry); Artificial intelligence; Computation; Natural language processing; Language model; Information retrieval; Theoretical computer science; Algorithm; Programming language","score_opus":0.034081906961631506,"score_gpt":0.29067125826414963,"score_spread":0.25658935130251814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2752083635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007918372,0.00019878669,0.9826477,0.00025106187,0.000060837225,0.00012140385,0.00056621706,0.005819629,0.0024159448],"genre_scores_gemma":[0.18212306,0.0002956486,0.80913776,0.00017932612,0.0002138079,0.0006908645,0.002954385,0.0007658566,0.003639257],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998059,0.00047353248,0.00014187237,0.00043653345,0.00068530743,0.00020377182],"domain_scores_gemma":[0.99692994,0.0012722946,0.00015979812,0.000895551,0.0006280782,0.000114324066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017363954,0.0011548332,0.0013310107,0.0022068846,0.0012502177,0.002629086,0.0037905644,0.00124008,0.011021453],"category_scores_gemma":[0.011601293,0.0006639457,0.0012981861,0.004358027,0.0009076431,0.0070777424,0.0029134012,0.0019011734,0.005374286],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005125925,0.0004083222,0.002425581,0.0004247799,0.00017647764,0.0003322569,0.00047709484,0.24430886,0.010428999,0.23467548,0.046551928,0.4592776],"study_design_scores_gemma":[0.000025552052,0.000025037934,0.00013322342,0.000006476153,0.0000131400675,0.000050589206,0.000038581456,0.9444264,0.0014874433,0.050171006,0.003607053,0.000015458027],"about_ca_topic_score_codex":0.009388079,"about_ca_topic_score_gemma":0.0122512765,"teacher_disagreement_score":0.011021453,"about_ca_system_score_codex":0.0015063701,"about_ca_system_score_gemma":0.0030839765,"threshold_uncertainty_score":0.03687048},"labels":[],"label_agreement":null},{"id":"W2752442988","doi":"10.1609/aaai.v32i1.11947","title":"Order-Planning Neural Text Generation From Structured Data","year":2018,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":101,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Computer science; Dialog box; Table (database); Natural language generation; Natural language processing; Artificial neural network; Artificial intelligence; Text generation; Encoder; NIST; Language model; Order (exchange); Component (thermodynamics); Natural language; Machine learning; Data mining","score_opus":0.22174273540355915,"score_gpt":0.34089337949952864,"score_spread":0.11915064409596948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2752442988","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12533163,0.0009799645,0.84802204,0.00083357736,0.0002724808,0.00042584466,0.0015503214,0.015740372,0.0068437965],"genre_scores_gemma":[0.6625438,0.00033881704,0.3215569,0.00040392537,0.00012842666,0.00045740453,0.004656025,0.0005155436,0.009399149],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995782,0.000106070685,0.000026718924,0.00017617008,0.00007455358,0.00003826083],"domain_scores_gemma":[0.99789643,0.0014054158,0.0000993765,0.00026932228,0.00025923233,0.00007028554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009595777,0.00090832484,0.00061360555,0.00086902623,0.00040270263,0.0006145742,0.001486357,0.00089599326,0.0034156493],"category_scores_gemma":[0.0042442787,0.0003642102,0.00067681837,0.00086815475,0.0004743912,0.0017524376,0.0007973885,0.0012383225,0.0010410313],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038363386,0.00035814187,0.001938088,0.00029077544,0.0000793118,0.00033421814,0.00030129545,0.40348703,0.01346808,0.009558996,0.017454252,0.55234617],"study_design_scores_gemma":[0.000026143449,0.000027095615,0.00013775231,0.000004600676,0.000009095009,0.000023714698,0.000017870176,0.989096,0.004076465,0.005788977,0.00078690815,0.0000054710617],"about_ca_topic_score_codex":0.0066353213,"about_ca_topic_score_gemma":0.012408413,"teacher_disagreement_score":0.0066353213,"about_ca_system_score_codex":0.0009327818,"about_ca_system_score_gemma":0.0008620327,"threshold_uncertainty_score":0.013193369},"labels":[],"label_agreement":null},{"id":"W2756560861","doi":"10.63317/2aovvs53yg45","title":"Second Order Co-occurrence PMI for Determining the Semantic Similarity of Words","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":109,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Semantic similarity; Natural language processing; Similarity (geometry); Computer science; Synonym (taxonomy); Artificial intelligence; Pointwise; Noun; Pointwise mutual information; Semantics (computer science); Test (biology); Linguistics; Mathematics; Mutual information","score_opus":0.03016811658200047,"score_gpt":0.28178269132258865,"score_spread":0.2516145747405882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2756560861","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056356315,0.0015144654,0.92985386,0.00016582807,0.00018660774,0.0006820117,0.0036621713,0.0028341645,0.004744565],"genre_scores_gemma":[0.384034,0.0006388266,0.605056,0.00005173698,0.00031614574,0.0019551923,0.0059144916,0.0004598898,0.0015736676],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935946,0.0015539348,0.0005052874,0.0016058596,0.00251035,0.00022995088],"domain_scores_gemma":[0.9864325,0.008654254,0.0010809149,0.0015588898,0.0019564573,0.00031698204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004386414,0.0015301694,0.0016457691,0.019021844,0.0014876825,0.0019562747,0.0019450474,0.0013024651,0.0028942025],"category_scores_gemma":[0.022153439,0.00070151826,0.001738323,0.012757031,0.0011263164,0.0042600124,0.0017487487,0.0017809583,0.0019860598],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011508032,0.0003910701,0.064720154,0.0013564294,0.0010387888,0.0005840396,0.0019749461,0.024535006,0.02158907,0.018843463,0.013349912,0.8504664],"study_design_scores_gemma":[0.00013633393,0.0003926988,0.081754155,0.00019812665,0.00044970572,0.0033750422,0.000886724,0.82963,0.020630954,0.039349407,0.022722853,0.00047397075],"about_ca_topic_score_codex":0.0043842676,"about_ca_topic_score_gemma":0.0047984994,"teacher_disagreement_score":0.019021844,"about_ca_system_score_codex":0.0014682638,"about_ca_system_score_gemma":0.0017308976,"threshold_uncertainty_score":0.02319783},"labels":[],"label_agreement":null},{"id":"W2757021967","doi":"10.1002/asi.23882","title":"Discourse relations in rationale‐containing text‐segments","year":2017,"lang":"en","type":"article","venue":"Journal of the Association for Information Science and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Generalizability theory; Computer science; Leverage (statistics); Perspective (graphical); Sample (material); Face (sociological concept); Empirical research; Discourse analysis; Face-to-face interaction; Data science; Artificial intelligence; Linguistics; Sociology; Psychology; Epistemology; Communication","score_opus":0.01801076566800524,"score_gpt":0.2963604448376717,"score_spread":0.27834967916966646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2757021967","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92772007,0.002384705,0.03940231,0.0008241249,0.00019707976,0.0006578077,0.014475318,0.0007033015,0.013635338],"genre_scores_gemma":[0.9268479,0.00065768615,0.047263518,0.00009935007,0.00011299659,0.00083452894,0.020749262,0.0001520083,0.0032827326],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99728787,0.0013017774,0.00026326426,0.0004939753,0.00053412915,0.000119018245],"domain_scores_gemma":[0.9671417,0.025171895,0.0037207804,0.0009911134,0.0026345463,0.0003398663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025065409,0.00051379926,0.00026626885,0.0053866906,0.001288318,0.0018001485,0.00037290118,0.0006852373,0.0030387251],"category_scores_gemma":[0.021192685,0.0001952012,0.00033361002,0.0047570146,0.0007458573,0.0026371214,0.0011030593,0.00077985716,0.0011186758],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024312881,0.0004582211,0.20084244,0.0065255044,0.00020415173,0.0028178259,0.17547691,0.0037374638,0.115222916,0.027811276,0.030257698,0.43421426],"study_design_scores_gemma":[0.00014063188,0.0005028408,0.50730723,0.0019252041,0.00027168143,0.0019294416,0.09359064,0.042309538,0.037569456,0.029426226,0.2847615,0.00026558517],"about_ca_topic_score_codex":0.0024639878,"about_ca_topic_score_gemma":0.0055227503,"teacher_disagreement_score":0.0053866906,"about_ca_system_score_codex":0.00095549884,"about_ca_system_score_gemma":0.00093606027,"threshold_uncertainty_score":0.013256013},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"},{"model":"gpt","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"agree"},{"id":"W2757442264","doi":"10.18653/v1/d17-1048","title":"A Cognition Based Attention Model for Sentiment Analysis","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Advanced Micro Devices (Canada)","funders":"Hong Kong Polytechnic University","keywords":"Computer science; Context (archaeology); Sentiment analysis; Cognition; Reading (process); Artificial intelligence; Eye tracking; Cognitive model; Natural language processing; Machine learning; Psychology; Linguistics","score_opus":0.06293940680468281,"score_gpt":0.306727946667495,"score_spread":0.24378853986281218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2757442264","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08300909,0.0012773586,0.90513283,0.0010768647,0.00027919371,0.00015748186,0.0006998076,0.0017730157,0.006594322],"genre_scores_gemma":[0.9024738,0.00063289254,0.08755281,0.00045834278,0.00026223925,0.00021899794,0.0007510186,0.000110110785,0.0075397594],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99966276,0.00006563905,0.000017899525,0.00014357026,0.000050500574,0.00005963818],"domain_scores_gemma":[0.99940884,0.00028970797,0.000059090926,0.00004333314,0.00016062887,0.00003849101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008538158,0.00093276624,0.0006441127,0.0011874518,0.00037591418,0.00091296335,0.0011331781,0.0009664954,0.0032568055],"category_scores_gemma":[0.0024268276,0.00034631815,0.0015654892,0.0009212821,0.00046242852,0.0014539909,0.00061621406,0.0014873786,0.0008665039],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065908575,0.0005631747,0.017003633,0.00036619193,0.0005679415,0.00042977155,0.0010456778,0.24190822,0.052504472,0.039141104,0.014429815,0.631381],"study_design_scores_gemma":[0.000013983006,0.00007025866,0.004514771,0.000015146455,0.00008574825,0.000057315454,0.000031285985,0.977629,0.0019308829,0.014246016,0.001384749,0.000020862037],"about_ca_topic_score_codex":0.013928413,"about_ca_topic_score_gemma":0.014019227,"teacher_disagreement_score":0.013928413,"about_ca_system_score_codex":0.0014189645,"about_ca_system_score_gemma":0.00074383226,"threshold_uncertainty_score":0.027694702},"labels":[],"label_agreement":null},{"id":"W2759162401","doi":"","title":"Named Entity Recognition and Hashtag Decomposition to Improve the Classification of Tweets","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec; Université du Québec à Montréal","funders":"","keywords":"WordNet; Computer science; Artificial intelligence; Natural language processing; Preprocessor; Named-entity recognition; Segmentation; Field (mathematics); Task (project management); Information retrieval; Semantics (computer science)","score_opus":0.0829395137147931,"score_gpt":0.3343411440966311,"score_spread":0.25140163038183805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2759162401","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24150814,0.0021721069,0.71530807,0.0013225849,0.0007874202,0.0006183474,0.0054628532,0.024133742,0.008686766],"genre_scores_gemma":[0.52411187,0.00075305905,0.45049724,0.00032342813,0.00030810368,0.00027492674,0.015440213,0.00047468327,0.007816541],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986726,0.00040273336,0.00017349534,0.0003228578,0.00027692132,0.00015151594],"domain_scores_gemma":[0.99658597,0.0013704859,0.0003719011,0.00052330375,0.0010177459,0.00013060105],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016448481,0.0010809359,0.0007878874,0.0055133067,0.00074044673,0.0017176085,0.00088476983,0.0010575756,0.0025558916],"category_scores_gemma":[0.004966146,0.00027762252,0.0010684483,0.0043633967,0.00033215908,0.004296219,0.00093681156,0.0010626945,0.0049005128],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001135274,0.0007552558,0.022905488,0.0005453112,0.00033289348,0.00049698714,0.00083550095,0.012818293,0.07755624,0.00754678,0.024395544,0.8506764],"study_design_scores_gemma":[0.00007785489,0.00036622124,0.03187007,0.00009507802,0.00036979557,0.00065509055,0.0010664214,0.7912063,0.1152668,0.017876668,0.040992647,0.0001571258],"about_ca_topic_score_codex":0.0036345446,"about_ca_topic_score_gemma":0.0050314707,"teacher_disagreement_score":0.0055133067,"about_ca_system_score_codex":0.0005720077,"about_ca_system_score_gemma":0.00078473665,"threshold_uncertainty_score":0.008698881},"labels":[],"label_agreement":null},{"id":"W2759598007","doi":"","title":"UQAM-NTL: Named entity recognition in Twitter messages.","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec; Université du Québec à Montréal","funders":"","keywords":"Conditional random field; Named-entity recognition; Computer science; Task (project management); Conjunction (astronomy); Artificial intelligence; Natural language processing; Named entity; Entity linking; Information retrieval; Machine learning; Knowledge base; Engineering","score_opus":0.0937972724176686,"score_gpt":0.3230069221583068,"score_spread":0.22920964974063818,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2759598007","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023659162,0.0009198162,0.14075145,0.0012572564,0.0007275176,0.0012287003,0.25287557,0.5577907,0.02078977],"genre_scores_gemma":[0.13361219,0.00065270445,0.2727251,0.00080338353,0.00031302994,0.0022193932,0.5464757,0.00931306,0.033885468],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998669,0.00032943234,0.00015569862,0.00036536093,0.00036530764,0.00011516685],"domain_scores_gemma":[0.9978136,0.0006607587,0.0002652151,0.0006858767,0.000425902,0.0001485513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017530058,0.0016057907,0.0010465836,0.0035284194,0.00090121705,0.001548775,0.001795361,0.0014209801,0.02592803],"category_scores_gemma":[0.008590053,0.0004967052,0.00057394203,0.001839298,0.00039004313,0.0048878123,0.0026888265,0.0009231648,0.042956788],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013451176,0.00027904275,0.010294253,0.0015163742,0.000162502,0.00093008974,0.00075454405,0.0040969523,0.031902086,0.004718136,0.6202483,0.32375252],"study_design_scores_gemma":[0.00030129417,0.00041835796,0.013321496,0.00033265154,0.00013224724,0.001031351,0.00070850167,0.28130808,0.1282512,0.01159046,0.5623227,0.00028165133],"about_ca_topic_score_codex":0.0077865543,"about_ca_topic_score_gemma":0.008304489,"teacher_disagreement_score":0.02592803,"about_ca_system_score_codex":0.0009729012,"about_ca_system_score_gemma":0.0011455384,"threshold_uncertainty_score":0.08673787},"labels":[],"label_agreement":null},{"id":"W2759604924","doi":"10.18653/v1/d17-1031","title":"Word Embeddings based on Fixed-Size Ordinally Forgetting Encoding","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Word (group theory); Computer science; Context (archaeology); Encoding (memory); ENCODE; Artificial intelligence; Dimension (graph theory); Word embedding; Similarity (geometry); Natural language processing; Speech recognition; Embedding; Algorithm; Pattern recognition (psychology); Mathematics; Combinatorics","score_opus":0.02546458904238419,"score_gpt":0.27697213808021137,"score_spread":0.2515075490378272,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2759604924","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03187282,0.00083206815,0.9639875,0.00015035177,0.00019492142,0.000055299726,0.00041666956,0.0017592721,0.00073113106],"genre_scores_gemma":[0.5333568,0.0013725415,0.45548844,0.00029940563,0.00029991788,0.00031136288,0.0033656582,0.0003749689,0.005130973],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994523,0.00013017945,0.000051848863,0.00018669124,0.00012376675,0.00005524159],"domain_scores_gemma":[0.9983551,0.0006328805,0.0001686813,0.00040752906,0.00035583065,0.000079961515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005598307,0.0013038521,0.000861854,0.0011579236,0.00027648563,0.00081798155,0.0013669281,0.0007198827,0.0025486508],"category_scores_gemma":[0.005379863,0.00038463477,0.00075031543,0.0014868212,0.00056741654,0.0039490513,0.0010053301,0.0016208714,0.0014390375],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032851525,0.00019848875,0.0029722531,0.00034173444,0.00011335808,0.00016463331,0.00028950258,0.06434383,0.017489225,0.018562166,0.0067067333,0.8884895],"study_design_scores_gemma":[0.00003404497,0.00021335552,0.00093017076,0.000048361915,0.000065754575,0.00022931657,0.00008779723,0.94856167,0.0094098775,0.035576276,0.0047921743,0.000051155097],"about_ca_topic_score_codex":0.002331127,"about_ca_topic_score_gemma":0.003951211,"teacher_disagreement_score":0.0025486508,"about_ca_system_score_codex":0.00039135435,"about_ca_system_score_gemma":0.0006591452,"threshold_uncertainty_score":0.008526146},"labels":[],"label_agreement":null},{"id":"W2763886647","doi":"10.18653/v1/w17-5535","title":"Exploring Joint Neural Model for Sentence Level Discourse Parsing and Sentiment Analysis","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Parsing; Sentence; Natural language processing; Sentiment analysis; Artificial intelligence; Task (project management); Joint (building); Set (abstract data type); Representation (politics); Relation (database); Artificial neural network; Data mining","score_opus":0.4377446350894824,"score_gpt":0.3489493024477397,"score_spread":0.08879533264174272,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2763886647","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1902826,0.0013395943,0.79897636,0.0009986527,0.00016225508,0.00012198558,0.00048681442,0.003240704,0.0043910933],"genre_scores_gemma":[0.81872106,0.00049037294,0.17114572,0.00026117242,0.0001178795,0.00025221772,0.0013249359,0.00021534534,0.007471243],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994931,0.00019450193,0.00002425962,0.0001669232,0.000054612578,0.00006656686],"domain_scores_gemma":[0.99872607,0.00078857393,0.00008269698,0.00009194974,0.00025372964,0.000056882895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018494494,0.0009396471,0.00070014474,0.0006842574,0.000377064,0.0011358535,0.0012977269,0.0012557195,0.0021370153],"category_scores_gemma":[0.0032572886,0.0005025019,0.0009049247,0.0006471355,0.000454492,0.0029106806,0.0009313284,0.0019880165,0.0007937064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062431005,0.00049418677,0.0035952714,0.0002064128,0.00023793262,0.00021632535,0.00036717224,0.6658664,0.016666625,0.012239294,0.004545651,0.2949404],"study_design_scores_gemma":[0.000004442578,0.000017829452,0.00018836494,0.0000034997322,0.000011782717,0.0000046926953,0.000009367869,0.99648356,0.0008064752,0.0022822253,0.000184238,0.0000034986174],"about_ca_topic_score_codex":0.009495934,"about_ca_topic_score_gemma":0.012650475,"teacher_disagreement_score":0.009495934,"about_ca_system_score_codex":0.00113761,"about_ca_system_score_gemma":0.0012871198,"threshold_uncertainty_score":0.018881321},"labels":[],"label_agreement":null},{"id":"W2766439719","doi":"","title":"Context effects in explanation evaluation - eScholarship","year":2014,"lang":"en","type":"article","venue":"Proceedings of the Annual Meeting of the Cognitive Science Society","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Counterintuitive; Narrative; Recall; Context (archaeology); Psychology; Plot (graphics); Cognition; Social psychology; Cognitive psychology; Epistemology; History; Literature; Art; Philosophy; Mathematics","score_opus":0.02109351497195494,"score_gpt":0.2759000787834067,"score_spread":0.2548065638114518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2766439719","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97061116,0.0023919488,0.0039793393,0.0010831149,0.00005227058,0.00012400991,0.00010499832,0.000114731,0.02153845],"genre_scores_gemma":[0.9974367,0.00025558055,0.0014939342,0.000051241594,0.000012331028,0.00003552944,0.00003501111,0.000019977859,0.00065962813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99430156,0.004211394,0.00023549313,0.0003785748,0.00067263335,0.00020035468],"domain_scores_gemma":[0.9231103,0.066714,0.004818402,0.0023748274,0.0018704038,0.0011120571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048628626,0.00027454758,0.00029320928,0.00047391962,0.0006734094,0.002007327,0.0005853379,0.0008213855,0.010411969],"category_scores_gemma":[0.069754615,0.0003290846,0.00038146295,0.0003542813,0.0007456647,0.003665836,0.001895138,0.00085624284,0.00042457652],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013010462,0.0035168284,0.22571573,0.0052249813,0.0010150963,0.0028295119,0.07193375,0.0063512474,0.033446625,0.033147477,0.006178351,0.59763],"study_design_scores_gemma":[0.001210084,0.005104363,0.83122313,0.0014768103,0.0014934592,0.0022022445,0.031888615,0.017148048,0.021527683,0.049849346,0.036593255,0.0002830062],"about_ca_topic_score_codex":0.0011454921,"about_ca_topic_score_gemma":0.001380515,"teacher_disagreement_score":0.010411969,"about_ca_system_score_codex":0.00090186426,"about_ca_system_score_gemma":0.0005486437,"threshold_uncertainty_score":0.034831524},"labels":[],"label_agreement":null},{"id":"W2767855269","doi":"","title":"A Connectionist Model of Semantic Memory: Superordinate structure without hierarchies","year":2001,"lang":"en","type":"article","venue":"eScholarship (California Digital Library)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Superordinate goals; Connectionism; Categorization; Combinatorics; Psychology; Computer science; Physics; Natural language processing; Artificial intelligence; Cognitive science; Cognitive psychology; Philosophy; Mathematics; Social psychology; Artificial neural network","score_opus":0.01748105153553088,"score_gpt":0.21461781444771422,"score_spread":0.19713676291218335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767855269","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.178876,0.0032115907,0.60352904,0.015366926,0.00037492928,0.00018400144,0.0008403942,0.00086348824,0.19675365],"genre_scores_gemma":[0.93606484,0.0011194816,0.046899375,0.0005858227,0.00026107224,0.00015010085,0.00030094516,0.00008098608,0.014537278],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994392,0.00017579184,0.000023735223,0.00016671595,0.000113981136,0.00008055724],"domain_scores_gemma":[0.99889356,0.000461977,0.00011733483,0.00021618464,0.00015122202,0.00015984388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001217051,0.00046994648,0.00062135287,0.0012526279,0.0008580707,0.0033510812,0.002700032,0.0018224118,0.008703013],"category_scores_gemma":[0.0028283952,0.0006964311,0.0011513464,0.0012034711,0.0033986163,0.012027187,0.0015555441,0.0020163974,0.0012465257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009419223,0.000044919012,0.0009652232,0.000075478856,0.00003639519,0.00014428883,0.00088614743,0.008785201,0.0008245082,0.9698854,0.0013647606,0.0168934],"study_design_scores_gemma":[0.00002610175,0.00002335559,0.0004509017,0.000012131717,0.0000139394915,0.00010793965,0.00006617673,0.035960518,0.00010274089,0.9616637,0.0015630351,0.00000937454],"about_ca_topic_score_codex":0.0052012787,"about_ca_topic_score_gemma":0.0035540403,"teacher_disagreement_score":0.008703013,"about_ca_system_score_codex":0.0023327654,"about_ca_system_score_gemma":0.001317562,"threshold_uncertainty_score":0.029114485},"labels":[],"label_agreement":null},{"id":"W2767929746","doi":"10.1145/3132847.3133138","title":"An Empirical Study of Embedding Features in Learning to Rank","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Toronto Metropolitan University","funders":"","keywords":"Embedding; Computer science; Rank (graph theory); Ranking (information retrieval); Learning to rank; Word embedding; Artificial intelligence; Word (group theory); Information retrieval; Empirical research; Machine learning; Mathematics; Statistics","score_opus":0.04669500883657094,"score_gpt":0.3860673460005514,"score_spread":0.33937233716398046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2767929746","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7468668,0.008440137,0.23258962,0.0019142237,0.00017804255,0.0002456111,0.0011764433,0.0005690297,0.008020057],"genre_scores_gemma":[0.9795832,0.00040391216,0.018475119,0.00007239495,0.000098694305,0.000050492366,0.0006893565,0.00005060237,0.00057622755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98277503,0.013052187,0.00057216984,0.0010183081,0.002160978,0.0004212994],"domain_scores_gemma":[0.67377335,0.29605666,0.009979675,0.012547381,0.006336783,0.0013062338],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022031596,0.0011173041,0.0011804568,0.002057259,0.0006762626,0.0018026618,0.0011178497,0.0013584509,0.0028970235],"category_scores_gemma":[0.18876345,0.00033347105,0.0006080806,0.003702571,0.0018858257,0.0063218754,0.0010307113,0.0026827445,0.0007088472],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028847253,0.0021685057,0.20499729,0.0013969688,0.0010040269,0.00033211138,0.0007755075,0.21254826,0.0027618865,0.035718232,0.011212737,0.5241998],"study_design_scores_gemma":[0.00019978547,0.0021136678,0.05028744,0.0001544771,0.0002185416,0.00069867715,0.00052297395,0.8973448,0.003807239,0.040483534,0.004025808,0.00014303441],"about_ca_topic_score_codex":0.0016854176,"about_ca_topic_score_gemma":0.0014164513,"teacher_disagreement_score":0.022031596,"about_ca_system_score_codex":0.00081291416,"about_ca_system_score_gemma":0.0005002906,"threshold_uncertainty_score":0.11651558},"labels":[],"label_agreement":null},{"id":"W2768091364","doi":"10.48550/arxiv.1711.02013","title":"Neural Language Modeling by Jointly Learning Syntax and Lexicon","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Parsing; Natural language processing; Syntax; Language model; Leverage (statistics); Tree structure; Artificial neural network; Lexicon; Recurrent neural network; Data structure; Programming language","score_opus":0.07486587234654331,"score_gpt":0.19691794624440412,"score_spread":0.12205207389786081,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2768091364","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021069264,0.00033319124,0.974076,0.00041818855,0.000057461537,0.000040223375,0.0003516631,0.0018646973,0.0017893056],"genre_scores_gemma":[0.6539182,0.00093988405,0.33051193,0.00044996804,0.0002037896,0.00040716564,0.0031285097,0.0006393463,0.00980124],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99959487,0.00013809596,0.000021508597,0.00014967275,0.00006020413,0.000035650075],"domain_scores_gemma":[0.99936444,0.00033761546,0.000085739455,0.00007892228,0.00010719762,0.000026082806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00069410336,0.0009872579,0.0007446148,0.0011546188,0.00028359875,0.0011004558,0.0018316794,0.00092196977,0.0017629286],"category_scores_gemma":[0.0027391047,0.0006565458,0.0012180292,0.00106577,0.0004969382,0.0035988463,0.0009547394,0.0017059523,0.0012290365],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014512852,0.00013677527,0.0018120913,0.00018068838,0.00022134594,0.00026107242,0.00020796155,0.65208393,0.010426771,0.048271537,0.0074390885,0.27881357],"study_design_scores_gemma":[0.000005623112,0.000010347113,0.00007019924,0.0000037470475,0.00001313261,0.000016759293,0.0000054181833,0.9790096,0.00064390316,0.019672582,0.00054224563,0.000006413851],"about_ca_topic_score_codex":0.0046792105,"about_ca_topic_score_gemma":0.008670865,"teacher_disagreement_score":0.0046792105,"about_ca_system_score_codex":0.00073450146,"about_ca_system_score_gemma":0.0011093868,"threshold_uncertainty_score":0.009303927},"labels":[],"label_agreement":null},{"id":"W2769778028","doi":"10.5555/3107979.3107984","title":"A comparative study on content-based paper-to-paper recommendation approaches in scientific literature","year":2017,"lang":"en","type":"article","venue":"Communications and Networking Symposium","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Information retrieval; Set (abstract data type); tf–idf; Word embedding; Word (group theory); Representation (politics); Domain (mathematical analysis); Term (time); Embedding; Recommender system; Data mining; Artificial intelligence; Mathematics","score_opus":0.23222467380403888,"score_gpt":0.33810592431048786,"score_spread":0.10588125050644898,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2769778028","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3750516,0.13879693,0.40554765,0.0035818985,0.0014666572,0.0024603112,0.0051701376,0.0075429436,0.06038193],"genre_scores_gemma":[0.6037229,0.024938526,0.35190162,0.0005619627,0.00089621975,0.0006619215,0.0068818545,0.00045036618,0.009984681],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9860643,0.004449531,0.0013209305,0.001361326,0.006302086,0.00050178],"domain_scores_gemma":[0.95442885,0.028182702,0.0021847496,0.003854237,0.010304331,0.0010450921],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.010021742,0.0012687402,0.0015106068,0.026024884,0.0011904355,0.0043287883,0.001803005,0.0021622179,0.0037375907],"category_scores_gemma":[0.04434249,0.0004475127,0.0023060322,0.022720361,0.0007394165,0.0058953627,0.0012054105,0.00089656975,0.0026700445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010072362,0.0007403871,0.02936801,0.0032186327,0.0015146927,0.00013569274,0.00091193325,0.0054349587,0.006537749,0.005101184,0.007507993,0.9385217],"study_design_scores_gemma":[0.00076534797,0.0067845625,0.28405812,0.0031601542,0.0070233583,0.004444936,0.008108634,0.42609435,0.050616264,0.026637238,0.18115519,0.0011518922],"about_ca_topic_score_codex":0.00593867,"about_ca_topic_score_gemma":0.007486577,"teacher_disagreement_score":0.9956712,"about_ca_system_score_codex":0.0016989408,"about_ca_system_score_gemma":0.0019527123,"threshold_uncertainty_score":0.05300069},"labels":[],"label_agreement":null},{"id":"W2770550542","doi":"10.1145/3278721.3278777","title":"Ethical Challenges in Data-Driven Dialogue Systems","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; McGill University","funders":"","keywords":"Offensive; Adversarial system; Conversation; Computer science; Process (computing); Reinforcement learning; Data science; Artificial intelligence; Psychology; Communication; Operations research; Engineering","score_opus":0.24997439749857267,"score_gpt":0.34192282597913476,"score_spread":0.09194842848056209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2770550542","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019621687,0.0038103394,0.8820804,0.070018336,0.0009006369,0.0003701478,0.00041301677,0.0005158024,0.02226965],"genre_scores_gemma":[0.67509973,0.0025591312,0.29901478,0.010993396,0.0016436727,0.0016135883,0.0007222197,0.0006475683,0.0077058924],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8601584,0.11181037,0.004509622,0.008013299,0.014110651,0.0013975567],"domain_scores_gemma":[0.6541605,0.27953944,0.008802317,0.037098467,0.017121227,0.0032780888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11207064,0.00085497514,0.0012567999,0.0013342737,0.004632613,0.011247196,0.0032363432,0.007064247,0.004484364],"category_scores_gemma":[0.27297586,0.0011854136,0.001011643,0.0011978175,0.01413231,0.014963201,0.010984683,0.00958994,0.0014954025],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010365318,0.000063430714,0.0019162981,0.000402721,0.000058073398,0.00025221938,0.0073914304,0.010514551,0.0012059942,0.92318815,0.0076008784,0.04730261],"study_design_scores_gemma":[0.000036535555,0.00003238711,0.00026966273,0.00020844638,0.000016322547,0.00022628538,0.0010083668,0.030192787,0.0013937609,0.9252736,0.041294426,0.00004730343],"about_ca_topic_score_codex":0.0014560478,"about_ca_topic_score_gemma":0.0009708456,"teacher_disagreement_score":0.11207064,"about_ca_system_score_codex":0.003810602,"about_ca_system_score_gemma":0.005657432,"threshold_uncertainty_score":0.5926933},"labels":[],"label_agreement":null},{"id":"W2773947245","doi":"","title":"MONPA: Multi-objective Named-entity and Part-of-speech Annotator for Chinese using Recurrent Neural Network","year":2017,"lang":"en","type":"article","venue":"International Joint Conference on Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Named-entity recognition; Sentence; Recurrent neural network; Task (project management); Segmentation; Speech recognition; Part of speech; Word (group theory); Artificial neural network; Named entity; Text segmentation","score_opus":0.05657633475324485,"score_gpt":0.35028846040631406,"score_spread":0.2937121256530692,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2773947245","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058952652,0.00084571843,0.8584614,0.00039326033,0.00032002802,0.00045286052,0.005649733,0.07050912,0.004415226],"genre_scores_gemma":[0.33424762,0.00042861243,0.61509806,0.0004161798,0.000110342466,0.00090224086,0.023752686,0.0028991178,0.022145184],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998833,0.00026414203,0.00005856735,0.0005419839,0.00017264209,0.00012961136],"domain_scores_gemma":[0.9985983,0.0004584061,0.00011521537,0.00030996744,0.00042801746,0.00009004408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029508406,0.0020527812,0.0011066719,0.0013783835,0.0012530073,0.0011906651,0.0026183284,0.001369752,0.0049904436],"category_scores_gemma":[0.0038896534,0.0007009138,0.0010863815,0.001143375,0.00048573007,0.002639902,0.0024140265,0.0014053817,0.0036925215],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013106938,0.00040657824,0.010049793,0.0008835337,0.0005902103,0.0012380516,0.00079198874,0.08983268,0.064682275,0.006075891,0.07937285,0.74476546],"study_design_scores_gemma":[0.000051939125,0.0001021906,0.00197625,0.00002550663,0.000096425756,0.0001293429,0.000111583824,0.96153873,0.023240266,0.0035782366,0.009084697,0.00006486975],"about_ca_topic_score_codex":0.03145958,"about_ca_topic_score_gemma":0.058632754,"teacher_disagreement_score":0.03145958,"about_ca_system_score_codex":0.0013043004,"about_ca_system_score_gemma":0.0031715382,"threshold_uncertainty_score":0.06255293},"labels":[],"label_agreement":null},{"id":"W2774599502","doi":"10.26615/978-954-452-049-6_069","title":"Recognizing Reputation Defence Strategies in Critical Political Exchanges","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; University of Pennsylvania","keywords":"Reputation; Argumentation theory; Computer science; Task (project management); Relation (database); Field (mathematics); Politics; Computer security; Artificial intelligence; Political science; Epistemology; Data mining; Management; Law; Mathematics; Economics","score_opus":0.12816844240752176,"score_gpt":0.38788495646436566,"score_spread":0.25971651405684393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2774599502","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56995714,0.0026355162,0.40599847,0.0022327767,0.00014843063,0.000334355,0.0014666962,0.002131995,0.015094633],"genre_scores_gemma":[0.9242647,0.00022051545,0.07162742,0.00014977048,0.00008619695,0.0000826858,0.0016647356,0.0000903369,0.0018137229],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955413,0.0020524058,0.0003440232,0.00096449297,0.0007774899,0.00032027447],"domain_scores_gemma":[0.9756276,0.014500944,0.004818692,0.0020636856,0.0022410878,0.0007480595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036932793,0.00074341445,0.0007300823,0.004886185,0.0012399678,0.0040953,0.001616139,0.0028948465,0.0014837373],"category_scores_gemma":[0.02431106,0.000464196,0.00088383333,0.002197575,0.0009766471,0.005294356,0.0019486488,0.0021770773,0.0011155731],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012996872,0.0009809618,0.18694001,0.0012034915,0.00045524698,0.0013603275,0.008945087,0.036612093,0.04919097,0.09731671,0.023085512,0.59260994],"study_design_scores_gemma":[0.000087128,0.00021838484,0.063440956,0.00021182495,0.00017270193,0.0013678307,0.0027227295,0.7417753,0.021113789,0.14382869,0.024931135,0.0001295705],"about_ca_topic_score_codex":0.0021999907,"about_ca_topic_score_gemma":0.0037898892,"teacher_disagreement_score":0.004886185,"about_ca_system_score_codex":0.0011473764,"about_ca_system_score_gemma":0.0008409254,"threshold_uncertainty_score":0.019532204},"labels":[],"label_agreement":null},{"id":"W2775337853","doi":"","title":"Assessing the Verifiability of Attributions in News Text","year":2017,"lang":"en","type":"article","venue":"International Joint Conference on Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Operationalization; Attribution; Verifiable secret sharing; Computer science; Rank (graph theory); Task (project management); Fidelity; Crowdsourcing; Statement (logic); Authorship attribution; Natural language processing; Information retrieval; Artificial intelligence; Psychology; Social psychology; World Wide Web; Linguistics; Set (abstract data type); Mathematics","score_opus":0.0672499933037815,"score_gpt":0.36785059382849045,"score_spread":0.30060060052470894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2775337853","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.872023,0.0012054458,0.1108798,0.0010425893,0.00031730175,0.00031500394,0.002546291,0.0008276884,0.010842915],"genre_scores_gemma":[0.9833265,0.00013598689,0.014309974,0.000056206027,0.00013472214,0.000067444664,0.0013345374,0.00006296337,0.0005716559],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9631476,0.017994696,0.0032712552,0.005432263,0.009135879,0.0010183585],"domain_scores_gemma":[0.42670295,0.4742189,0.05132763,0.023384642,0.02229226,0.0020736163],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.051419243,0.001244263,0.0010034422,0.009632406,0.001718577,0.0059009064,0.0018109616,0.0027889248,0.002760623],"category_scores_gemma":[0.32810107,0.0006121808,0.0009330604,0.006311119,0.0032032414,0.007972193,0.004017776,0.0024939778,0.0013271525],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027225213,0.00093326345,0.58027536,0.0020763527,0.0013595462,0.0014125332,0.0138393035,0.042315785,0.020909652,0.01610883,0.00735031,0.31069657],"study_design_scores_gemma":[0.00022599124,0.0008520853,0.42986313,0.0006541682,0.0005838742,0.0012637051,0.0062380936,0.4427297,0.043637984,0.059361156,0.013965103,0.00062500866],"about_ca_topic_score_codex":0.0050979652,"about_ca_topic_score_gemma":0.0044382443,"teacher_disagreement_score":0.051419243,"about_ca_system_score_codex":0.0015319699,"about_ca_system_score_gemma":0.0014130695,"threshold_uncertainty_score":0.2719342},"labels":[],"label_agreement":null},{"id":"W2775747321","doi":"","title":"WiNER: A Wikipedia Annotated Corpus for Named Entity Recognition","year":2017,"lang":"en","type":"article","venue":"International Joint Conference on Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Natural language processing; Information retrieval; Artificial intelligence; Quality (philosophy); Named-entity recognition; Simple (philosophy); Training set; Range (aeronautics); Task (project management)","score_opus":0.060187248013458684,"score_gpt":0.3217091971620433,"score_spread":0.2615219491485846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2775747321","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07785952,0.0030065887,0.25240153,0.0014517785,0.0022081533,0.0016927245,0.5728976,0.046183966,0.04229807],"genre_scores_gemma":[0.06207394,0.0006622871,0.21829925,0.00033622747,0.00023114007,0.0012813194,0.70595074,0.0022933905,0.008871692],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99783176,0.0005427693,0.00035039993,0.0006255205,0.00052587717,0.00012364316],"domain_scores_gemma":[0.99224234,0.002378907,0.00064644124,0.0015840787,0.0026592081,0.00048912206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019378924,0.0011560015,0.00076066627,0.008492736,0.0013319072,0.0012219516,0.0017273846,0.0012162577,0.008344822],"category_scores_gemma":[0.010094488,0.00063076866,0.00062842004,0.0052913683,0.00052764284,0.0033414438,0.0016657361,0.0016861679,0.0065296474],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005942069,0.000692555,0.010057348,0.0038305,0.00023978816,0.0018390528,0.0015869214,0.006727236,0.039537318,0.013696126,0.62518144,0.2960176],"study_design_scores_gemma":[0.00017721068,0.00023306771,0.020004632,0.00048007496,0.00016009573,0.0017651281,0.00074258517,0.03939936,0.03722933,0.008827311,0.89072305,0.00025813206],"about_ca_topic_score_codex":0.008997921,"about_ca_topic_score_gemma":0.018633133,"teacher_disagreement_score":0.008997921,"about_ca_system_score_codex":0.0005825823,"about_ca_system_score_gemma":0.0021042288,"threshold_uncertainty_score":0.027916253},"labels":[],"label_agreement":null},{"id":"W2777929561","doi":"10.48550/arxiv.1712.08207","title":"Variational Attention for Sequence-to-Sequence Models","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Sequence (biology); Encoder; Computer science; Artificial neural network; Set (abstract data type); Gaussian; Algorithm; Latent variable; Random sequence; Artificial intelligence; Mechanism (biology); Mathematics; Physics","score_opus":0.2599454264256302,"score_gpt":0.24269512061273302,"score_spread":0.01725030581289716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2777929561","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064155636,0.00058124267,0.9912845,0.0003558518,0.0000428379,0.000031539123,0.00012746491,0.00024234837,0.0009185756],"genre_scores_gemma":[0.670016,0.0016838647,0.31270057,0.00066369347,0.00037541855,0.00048568362,0.0013221357,0.00046336252,0.0122893015],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983022,0.00081695267,0.00007768128,0.00044862155,0.00022834195,0.00012613703],"domain_scores_gemma":[0.99501604,0.0041524586,0.00018023474,0.0002877491,0.00025349585,0.00011006534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033717926,0.0011166146,0.0012985213,0.0011011118,0.0006328974,0.0012673211,0.0025976123,0.0022233815,0.0038295991],"category_scores_gemma":[0.011828739,0.00093316,0.0013099686,0.0012182136,0.0012857395,0.0026548542,0.0020182722,0.002758711,0.0006886747],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012204274,0.00007077465,0.0008625399,0.00016268385,0.00012655604,0.00015076592,0.0002754378,0.73799795,0.0027969142,0.19361176,0.0037228423,0.06009982],"study_design_scores_gemma":[0.000005608513,0.000008512628,0.00007204954,0.0000044867766,0.0000066552334,0.000015318974,0.0000054503366,0.9545179,0.00019883046,0.044601265,0.00055785343,0.000005932732],"about_ca_topic_score_codex":0.010128767,"about_ca_topic_score_gemma":0.010349351,"teacher_disagreement_score":0.010128767,"about_ca_system_score_codex":0.0020424721,"about_ca_system_score_gemma":0.0014594112,"threshold_uncertainty_score":0.020139635},"labels":[],"label_agreement":null},{"id":"W2778745504","doi":"10.1007/978-3-319-77028-4_55","title":"Dual Long Short-Term Memory Networks for Sub-Character Representation Learning","year":2018,"lang":"en","type":"book-chapter","venue":"Advances in intelligent systems and computing","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Algoma University","funders":"","keywords":"Computer science; Character (mathematics); Natural language processing; Artificial intelligence; Embedding; Segmentation; Term (time); Chinese characters; Boosting (machine learning)","score_opus":0.03502591276767502,"score_gpt":0.2958630473810277,"score_spread":0.2608371346133527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2778745504","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028522542,0.009890043,0.9387871,0.0009139702,0.0008965304,0.000053949083,0.00073073973,0.0027914634,0.017413588],"genre_scores_gemma":[0.5243767,0.007918153,0.37702894,0.0006598218,0.0008299878,0.0002044055,0.0034478584,0.00049591355,0.0850381],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99987185,0.000019438894,0.0000082502465,0.00005398957,0.0000252342,0.000021172398],"domain_scores_gemma":[0.9997279,0.00009323963,0.000019834053,0.000058317604,0.000075241995,0.000025483685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00031426264,0.00075838924,0.0005278214,0.0005649337,0.00027358913,0.0010090504,0.0014282499,0.0009654233,0.008089325],"category_scores_gemma":[0.0009511649,0.0002409388,0.0005191388,0.00091509655,0.00026660724,0.0019166579,0.00095629215,0.0018209915,0.0037640939],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017754608,0.00011200747,0.00037328876,0.00019488449,0.00007543449,0.000066367786,0.00005170533,0.03591426,0.015852401,0.017876856,0.016814828,0.91249037],"study_design_scores_gemma":[0.000011521186,0.00007564326,0.00036071741,0.000038859238,0.000058713154,0.00009955804,0.000032276515,0.93397266,0.010904093,0.040681098,0.013745551,0.000019287549],"about_ca_topic_score_codex":0.0018138654,"about_ca_topic_score_gemma":0.0031001642,"teacher_disagreement_score":0.008089325,"about_ca_system_score_codex":0.00052943407,"about_ca_system_score_gemma":0.00042650598,"threshold_uncertainty_score":0.027061522},"labels":[],"label_agreement":null},{"id":"W2779944493","doi":"10.24124/2010/bpgub700","title":"Study of document retrieval using Latent Semantic Indexing (LSI) on a very large data set.","year":2010,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Northern British Columbia","keywords":"Information retrieval; Computer science; Set (abstract data type); Search engine indexing; Programming language","score_opus":0.08128817225861587,"score_gpt":0.3508740158644802,"score_spread":0.26958584360586435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2779944493","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82055324,0.029615635,0.1316401,0.004632481,0.0003452268,0.000577576,0.0059479317,0.0016585992,0.005029146],"genre_scores_gemma":[0.9079726,0.0026987456,0.07699034,0.00018217148,0.00049593707,0.00029507795,0.0090408595,0.00017455786,0.0021497395],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99397403,0.0035706428,0.0005284297,0.00064740423,0.0010435962,0.0002359424],"domain_scores_gemma":[0.9104099,0.07967669,0.002017936,0.004671729,0.002529014,0.0006947258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012051525,0.0006341049,0.0018012064,0.005268139,0.00092048856,0.0024727443,0.001476885,0.0010276937,0.0014696557],"category_scores_gemma":[0.05073572,0.00036379864,0.0013006985,0.0053492174,0.0011750534,0.005602155,0.0009720043,0.0014497683,0.0006066565],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00386327,0.004845774,0.09517097,0.003926185,0.0025591846,0.0012071002,0.0022931066,0.14560074,0.014939471,0.041625146,0.05889675,0.6250723],"study_design_scores_gemma":[0.00017474203,0.0007051301,0.020561052,0.000072975294,0.00023371521,0.00069657725,0.0007069471,0.95885456,0.0036841836,0.009905908,0.004343953,0.00006028831],"about_ca_topic_score_codex":0.0075914683,"about_ca_topic_score_gemma":0.007944163,"teacher_disagreement_score":0.012051525,"about_ca_system_score_codex":0.0018201129,"about_ca_system_score_gemma":0.0015121347,"threshold_uncertainty_score":0.063735306},"labels":[],"label_agreement":null},{"id":"W2780329366","doi":"10.1142/s1793351x17400207","title":"On the Influence of Contextual Features for the Identification of Complex Words","year":2017,"lang":"en","type":"article","venue":"International Journal of Semantic Computing","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Medical Research Council; Natural Sciences and Engineering Research Council of Canada; Medical Research Council Canada","keywords":"Computer science; SemEval; Word (group theory); Artificial intelligence; Natural language processing; Identification (biology); Context (archaeology); Set (abstract data type); Natural language; Word identification; Word recognition; Linguistics; Task (project management)","score_opus":0.043467681168216615,"score_gpt":0.3322063310706441,"score_spread":0.2887386499024275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2780329366","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89640296,0.0067516486,0.08881075,0.0005317125,0.00025617436,0.00023841418,0.00032941488,0.001194851,0.005484154],"genre_scores_gemma":[0.97276956,0.0005850303,0.02532091,0.00011006282,0.00014796003,0.000038269194,0.00049427454,0.00015618649,0.00037781763],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935129,0.00351697,0.0004083495,0.0015065401,0.00074627384,0.0003089936],"domain_scores_gemma":[0.91032565,0.08064332,0.0018324325,0.0027617486,0.0033853045,0.0010516219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011039857,0.0017430824,0.0011404331,0.0020300818,0.0012515712,0.0023439517,0.0006420558,0.0009831768,0.0013545493],"category_scores_gemma":[0.052316315,0.0004907266,0.0009991198,0.000979299,0.0014947341,0.0045784563,0.0016469683,0.002129218,0.000650072],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009663403,0.0011492688,0.115661666,0.0012033624,0.0013620151,0.00048447086,0.0017490212,0.08245043,0.096956685,0.001958101,0.0035271172,0.68383443],"study_design_scores_gemma":[0.00024503132,0.00409237,0.11376937,0.0003172335,0.0019381106,0.0008051038,0.0012346572,0.7931672,0.07295068,0.0066508697,0.0045050434,0.00032444258],"about_ca_topic_score_codex":0.004675698,"about_ca_topic_score_gemma":0.008398733,"teacher_disagreement_score":0.011039857,"about_ca_system_score_codex":0.00056244823,"about_ca_system_score_gemma":0.0007957621,"threshold_uncertainty_score":0.058385015},"labels":[],"label_agreement":null},{"id":"W2781014390","doi":"10.1145/3086512.3086550","title":"Two-step cascaded textual entailment for legal bar exam question answering","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Institute of Informatics; Alberta Machine Intelligence Institute","keywords":"Textual entailment; Computer science; Logical consequence; Natural language processing; Artificial intelligence; Question answering; Negation; Information retrieval; Representation (politics); Programming language","score_opus":0.035222153940837955,"score_gpt":0.30926485816787647,"score_spread":0.2740427042270385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2781014390","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08829134,0.0012616502,0.8666646,0.0009872236,0.00014863597,0.0016997985,0.00480516,0.029089412,0.0070521133],"genre_scores_gemma":[0.301654,0.0003680539,0.6731078,0.00044064427,0.00017934445,0.000591613,0.017349252,0.000498847,0.005810486],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99758005,0.0005692553,0.00025563262,0.0007060825,0.0007129086,0.00017596455],"domain_scores_gemma":[0.99758255,0.0012672899,0.00015686676,0.00028038307,0.0005885941,0.00012438223],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017856072,0.0013943333,0.0009306663,0.00393179,0.00092875597,0.0013092428,0.0021720706,0.0014584201,0.009419363],"category_scores_gemma":[0.006913967,0.00051241537,0.0019751678,0.001643706,0.00049675157,0.0043471055,0.0021417742,0.0016542532,0.0038510715],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009548871,0.0010875044,0.0044749556,0.0011701725,0.00025465124,0.0008589515,0.001154224,0.016247315,0.077778876,0.014909885,0.029753186,0.8513553],"study_design_scores_gemma":[0.00016372475,0.0003461567,0.0044458234,0.00006915901,0.000276597,0.00083409844,0.00037384746,0.87624115,0.0694732,0.022501418,0.025172897,0.00010195154],"about_ca_topic_score_codex":0.010810379,"about_ca_topic_score_gemma":0.016780997,"teacher_disagreement_score":0.010810379,"about_ca_system_score_codex":0.0014675313,"about_ca_system_score_gemma":0.0018328697,"threshold_uncertainty_score":0.03151089},"labels":[],"label_agreement":null},{"id":"W2781697036","doi":"10.1007/978-3-319-73706-5_13","title":"In-Memory Distributed Training of Linear-Chain Conditional Random Fields with an Application to Fine-Grained Named Entity Recognition","year":2018,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Computer science; Conditional random field; Scalability; Named-entity recognition; Task (project management); Sequence labeling; Feature (linguistics); Sequence (biology); Artificial intelligence; Pattern recognition (psychology); Database","score_opus":0.023987875617359336,"score_gpt":0.25610979584764443,"score_spread":0.2321219202302851,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2781697036","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03974031,0.00022181594,0.94548565,0.00022062249,0.0000717799,0.00007345703,0.00017314282,0.0128257545,0.0011874493],"genre_scores_gemma":[0.4615588,0.00014623531,0.53249127,0.0001504751,0.00007333221,0.0001861958,0.00083968823,0.0007024497,0.0038514957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995474,0.00011595647,0.000026786578,0.0001775846,0.000079837664,0.000052335203],"domain_scores_gemma":[0.9981365,0.0010246485,0.000079709615,0.0004496105,0.00023586747,0.00007366388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014683456,0.0008268696,0.000888187,0.0004392603,0.0004925531,0.00074461836,0.0016763166,0.0010412025,0.004261066],"category_scores_gemma":[0.0046666926,0.00045994952,0.00044576923,0.00078309915,0.00051402504,0.0016957105,0.0010865736,0.0014089717,0.0015306212],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008756706,0.00034543689,0.0027069522,0.000117111136,0.00010604344,0.00026501014,0.0002573039,0.4576385,0.014028324,0.004412583,0.010865615,0.50838155],"study_design_scores_gemma":[0.000027151134,0.000019747502,0.00016062056,0.0000023301586,0.0000046496207,0.000020189022,0.0000106373645,0.99310195,0.003374077,0.0026848393,0.0005895974,0.0000042162505],"about_ca_topic_score_codex":0.009224933,"about_ca_topic_score_gemma":0.012829793,"teacher_disagreement_score":0.009224933,"about_ca_system_score_codex":0.0008124715,"about_ca_system_score_gemma":0.0009819766,"threshold_uncertainty_score":0.018342435},"labels":[],"label_agreement":null},{"id":"W2781827132","doi":"10.2196/medinform.8751","title":"A Pilot Study of Biomedical Text Comprehension using an Attention-Based Deep Neural Reader: Design and Experimental Analysis","year":2018,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Research Foundation of Korea; Ministry of Science, ICT and Future Planning; National Research Foundation","keywords":"Comprehension; Computer science; Artificial intelligence; Natural language processing; Reading comprehension; Context (archaeology); Task (project management); Deep learning; Domain (mathematical analysis); Process (computing); Reading (process); Linguistics; Engineering","score_opus":0.08271124705486638,"score_gpt":0.34488075539869206,"score_spread":0.2621695083438257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2781827132","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9515513,0.00039743024,0.03162839,0.00026794273,0.00018374542,0.011547474,0.0015859709,0.0006309576,0.0022068224],"genre_scores_gemma":[0.8869177,0.00045050879,0.07015765,0.0006775411,0.00015925374,0.032978576,0.0026317236,0.000153363,0.0058737653],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99878114,0.0005315426,0.00013489478,0.00027245548,0.00016350563,0.00011648614],"domain_scores_gemma":[0.99186957,0.005082305,0.00048738855,0.0007708618,0.0013463687,0.00044342838],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054731537,0.0011904025,0.0009073171,0.0005391835,0.00037095693,0.0006846844,0.0016590245,0.0014736953,0.0070580556],"category_scores_gemma":[0.009847118,0.00060010015,0.0009980182,0.00037577463,0.00085420883,0.0012178927,0.00078341935,0.001502903,0.0017080138],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.048708152,0.17614274,0.03258402,0.009587198,0.0018406939,0.0021789775,0.0037110287,0.068862416,0.17348236,0.0050404333,0.01722662,0.46063542],"study_design_scores_gemma":[0.029856008,0.3890135,0.087520435,0.00056053925,0.0028072651,0.0013862555,0.002042208,0.31718796,0.13126752,0.012686695,0.025056211,0.0006153922],"about_ca_topic_score_codex":0.0018159343,"about_ca_topic_score_gemma":0.0018329595,"teacher_disagreement_score":0.0070580556,"about_ca_system_score_codex":0.00095404155,"about_ca_system_score_gemma":0.0010633451,"threshold_uncertainty_score":0.028945148},"labels":[],"label_agreement":null},{"id":"W2782842435","doi":"","title":"Federating natural language question answering services of a cognitive enterprise data platform","year":2017,"lang":"en","type":"article","venue":"Computer Science and Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Computer science; Scalability; Question answering; Ranking (information retrieval); Natural language; Information retrieval; World Wide Web; Service (business); Artificial intelligence; Data science; Database","score_opus":0.015405816802963431,"score_gpt":0.26515971328529364,"score_spread":0.2497538964823302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2782842435","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23591901,0.0001768805,0.70818263,0.002521512,0.00016815797,0.00096000195,0.0011632049,0.04160652,0.009302112],"genre_scores_gemma":[0.61660093,0.000090760645,0.3722869,0.00095764437,0.000062385836,0.000476109,0.003331482,0.00069449615,0.005499208],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747616,0.00057774776,0.00017561598,0.0006904759,0.00076296396,0.00031704316],"domain_scores_gemma":[0.994311,0.0021785444,0.00020629652,0.0015332848,0.0013155838,0.00045531266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042653456,0.0008517336,0.00072809623,0.0013154836,0.0008410152,0.0026062054,0.0024688095,0.001629322,0.0029748897],"category_scores_gemma":[0.017808525,0.00063482206,0.0013424117,0.0009745627,0.0011996466,0.0045858333,0.004142589,0.0027597737,0.0014299558],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001639415,0.0020632292,0.024691971,0.00045688014,0.00037763207,0.0021955683,0.003251923,0.21041678,0.06556568,0.07174915,0.025005486,0.5925862],"study_design_scores_gemma":[0.00005327679,0.0001190768,0.0018258408,0.000024714085,0.000044809593,0.00011926113,0.00031552406,0.93116957,0.021032952,0.03573483,0.009499554,0.000060499384],"about_ca_topic_score_codex":0.016678203,"about_ca_topic_score_gemma":0.012716786,"teacher_disagreement_score":0.016678203,"about_ca_system_score_codex":0.0022221727,"about_ca_system_score_gemma":0.0032472627,"threshold_uncertainty_score":0.033162296},"labels":[],"label_agreement":null},{"id":"W2783530388","doi":"10.1111/coin.12152","title":"Improving text relatedness by incorporating phrase relatedness with word relatedness","year":2018,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"SemEval; Phrase; Computer science; Word (group theory); Natural language processing; Artificial intelligence; n-gram; Mathematics; Task (project management); Language model","score_opus":0.02036011953862601,"score_gpt":0.25799002373788116,"score_spread":0.23762990419925514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783530388","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5617146,0.010032024,0.38469338,0.0011018408,0.0006481095,0.0006553603,0.0075907074,0.011651044,0.021913012],"genre_scores_gemma":[0.8446369,0.00097111234,0.13207862,0.0004046145,0.00040608965,0.00028840732,0.014990868,0.0006088186,0.0056145815],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959223,0.001697334,0.00029521933,0.0009900136,0.00086405297,0.00023096334],"domain_scores_gemma":[0.9952003,0.002196483,0.00057983893,0.000673598,0.0011034034,0.00024639122],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029597962,0.0018621136,0.0011263223,0.0064196656,0.0009366076,0.0015605289,0.0010183171,0.0015777707,0.0028773164],"category_scores_gemma":[0.014072119,0.0003207808,0.0013079105,0.0035109012,0.000638184,0.0050265295,0.0026660655,0.0013650546,0.0033525499],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015917914,0.0013546204,0.06294082,0.0018623262,0.00095369737,0.00069950946,0.0013931163,0.059921004,0.052236844,0.008026171,0.05658326,0.7524368],"study_design_scores_gemma":[0.00021065524,0.0018620034,0.060289543,0.0003100938,0.0007697744,0.0011647842,0.0011052187,0.8261307,0.034162223,0.03694007,0.036791734,0.00026317846],"about_ca_topic_score_codex":0.0035911805,"about_ca_topic_score_gemma":0.005251914,"teacher_disagreement_score":0.0064196656,"about_ca_system_score_codex":0.0006013434,"about_ca_system_score_gemma":0.00084052567,"threshold_uncertainty_score":0.015653133},"labels":[],"label_agreement":null},{"id":"W2784152338","doi":"10.1109/icmla.2017.0-122","title":"Cybersecurity Automated Information Extraction Techniques: Drawbacks of Current Methods, and Enhanced Extractors","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Science and Technology Directorate; Ministère de la Défense Nationale; UT-Battelle; Battelle; U.S. Department of Homeland Security; U.S. Department of Energy","keywords":"Overfitting; Computer science; Leverage (statistics); Information extraction; Relationship extraction; Identification (biology); Information retrieval; Domain (mathematical analysis); Precision and recall; Recall; Artificial intelligence; Natural language processing; Data mining","score_opus":0.025526095890904904,"score_gpt":0.40030891763126164,"score_spread":0.3747828217403567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2784152338","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04309054,0.014699389,0.9197701,0.0045977603,0.00037041324,0.00036097094,0.001622137,0.008851316,0.006637445],"genre_scores_gemma":[0.1889296,0.012109783,0.7815891,0.0010973846,0.0007574336,0.0004510715,0.0054810494,0.0011811792,0.008403362],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9878293,0.0049113086,0.0009521937,0.0023873844,0.003603932,0.00031591122],"domain_scores_gemma":[0.9455993,0.032849453,0.002227677,0.012625894,0.0063188733,0.0003788474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016836964,0.002779928,0.001968925,0.0085115405,0.0015543737,0.006449224,0.0034958425,0.0019403618,0.002726944],"category_scores_gemma":[0.03664212,0.0014967738,0.0017358036,0.0067306454,0.0020901663,0.0134193,0.003870948,0.0033338289,0.0058956807],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022564361,0.00033911678,0.008126525,0.0018151146,0.00036289968,0.00017231489,0.0013368796,0.014650294,0.010834632,0.009419341,0.012445256,0.94027203],"study_design_scores_gemma":[0.00014545271,0.0007004191,0.015793575,0.0021599277,0.0008112755,0.0027175923,0.0026405875,0.63844764,0.08006972,0.079814166,0.17621939,0.00048028337],"about_ca_topic_score_codex":0.004174157,"about_ca_topic_score_gemma":0.006395507,"teacher_disagreement_score":0.016836964,"about_ca_system_score_codex":0.0011031801,"about_ca_system_score_gemma":0.0024687531,"threshold_uncertainty_score":0.08904338},"labels":[],"label_agreement":null},{"id":"W2785517012","doi":"10.2196/medinform.8662","title":"Automated Information Extraction on Treatment and Prognosis for Non–Small Cell Lung Cancer Radiotherapy Patients: Clinical Study","year":2018,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Medicine; Data extraction; Lung cancer; Radiation therapy; Information extraction; Recall; Medical physics; Computer science; Oncology; Internal medicine; MEDLINE; Information retrieval","score_opus":0.0318367729255427,"score_gpt":0.3748708232248038,"score_spread":0.3430340502992611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2785517012","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9019069,0.0052317623,0.058060203,0.0009396605,0.00010563693,0.001982623,0.025084166,0.0025869007,0.0041021355],"genre_scores_gemma":[0.87338775,0.0015884326,0.08518099,0.00021277479,0.000086254244,0.0010922027,0.037311807,0.00009777634,0.0010420568],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9974076,0.00095196656,0.0005743719,0.0005768642,0.00039727625,0.00009184517],"domain_scores_gemma":[0.9896366,0.006141989,0.0011526849,0.00096913555,0.0019348358,0.00016466522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046099224,0.00052080315,0.0005613958,0.0037402252,0.00035982154,0.0009246794,0.00053043745,0.00054431055,0.00232788],"category_scores_gemma":[0.016133908,0.00016372334,0.00089399784,0.0023180433,0.00034356012,0.000890477,0.00091752986,0.00031320978,0.00081360724],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002059363,0.0007086383,0.348538,0.0030431102,0.00046985812,0.001254885,0.0013375909,0.0059725563,0.012610328,0.0003643304,0.008299965,0.6153414],"study_design_scores_gemma":[0.001047595,0.0033729991,0.78347933,0.0011545875,0.0027930671,0.0056527876,0.0025594423,0.07727933,0.06854527,0.0028857766,0.050987728,0.00024205162],"about_ca_topic_score_codex":0.0017250058,"about_ca_topic_score_gemma":0.0017526353,"teacher_disagreement_score":0.0046099224,"about_ca_system_score_codex":0.00044277188,"about_ca_system_score_gemma":0.0011820378,"threshold_uncertainty_score":0.024379909},"labels":[],"label_agreement":null},{"id":"W2785543907","doi":"10.24963/ijcai.2018/631","title":"Generating Thematic Chinese Poetry using Conditional Variational Autoencoders with Hybrid Decoders","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Ningbo Municipal Science and Technology Innovative Research Team; Natural Sciences and Engineering Research Council of Canada; National Key Research and Development Program of China; Beijing Advanced Innovation Center for Imaging Technology","keywords":"Autoencoder; Word2vec; Computer science; Poetry; Context (archaeology); Artificial intelligence; Natural language processing; Artificial neural network; Relevance (law); Sequence (biology); Theme (computing); Recurrent neural network; Linguistics","score_opus":0.029440764706389528,"score_gpt":0.27853909504118407,"score_spread":0.24909833033479453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2785543907","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043182373,0.00021384226,0.95273155,0.0002049992,0.00006449516,0.00004288379,0.00012111979,0.0010747081,0.0023640594],"genre_scores_gemma":[0.69632345,0.00025926356,0.29352304,0.00021546612,0.00007181051,0.00014747787,0.0008057961,0.0002923241,0.0083614355],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999744,0.00008732737,0.000014446283,0.000074269854,0.000050610222,0.00002928511],"domain_scores_gemma":[0.99939764,0.0003611387,0.00003184241,0.00006546931,0.000115106646,0.0000288272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00062033365,0.0006021515,0.00047194236,0.00030262352,0.00020546305,0.0004945823,0.00069682836,0.0005786173,0.0020653866],"category_scores_gemma":[0.0020054223,0.00036626417,0.0006604663,0.0003052753,0.00041175526,0.0008026004,0.00072588853,0.0009902172,0.0007487929],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016961027,0.0001365154,0.0013582564,0.00015651,0.00011169202,0.00016830029,0.00020862927,0.59165466,0.022241233,0.020391386,0.003896957,0.35950628],"study_design_scores_gemma":[0.0000052657138,0.000011407649,0.000062127016,0.0000027344956,0.000004537753,0.000012301693,0.00000704754,0.9955479,0.002178695,0.0018406792,0.0003245935,0.000002760546],"about_ca_topic_score_codex":0.0027804316,"about_ca_topic_score_gemma":0.0038569919,"teacher_disagreement_score":0.0027804316,"about_ca_system_score_codex":0.00034348268,"about_ca_system_score_gemma":0.0005696994,"threshold_uncertainty_score":0.00690943},"labels":[],"label_agreement":null},{"id":"W2786028756","doi":"10.3166/ria.31.619-648","title":"Amélioration continue d’une chaîne de traitement de documents avec l’apprentissage par renforcement","year":2017,"lang":"fr","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Humanities; Political science; Computer science; Philosophy","score_opus":0.0630402083579833,"score_gpt":0.30841906095344407,"score_spread":0.24537885259546077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786028756","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10663319,0.0013676403,0.87955195,0.0007888113,0.00019332515,0.0002106738,0.0003105263,0.004713405,0.006230399],"genre_scores_gemma":[0.67872953,0.0011209535,0.29919443,0.00022454224,0.00010276449,0.0003355417,0.0005492625,0.00084044604,0.018902456],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965861,0.00072649453,0.00024777016,0.0008257297,0.0013517056,0.00026215406],"domain_scores_gemma":[0.98703927,0.0066248,0.000939691,0.0019103064,0.0029725337,0.0005133769],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036778154,0.0012237787,0.000981494,0.0015679552,0.0011650982,0.004068244,0.0014042257,0.0022043379,0.0065929005],"category_scores_gemma":[0.01816599,0.0007612935,0.0013974054,0.0010795842,0.0012139325,0.004344143,0.0019228167,0.002131146,0.002163738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010439823,0.00036096378,0.008782905,0.00088031177,0.00020131175,0.00074609276,0.0020550087,0.51782393,0.08668093,0.027415149,0.0033585418,0.35065088],"study_design_scores_gemma":[0.00008166772,0.0006240485,0.0033756194,0.00014234995,0.00012889985,0.000350337,0.00037758806,0.89845306,0.06025008,0.015232655,0.020864978,0.000118654374],"about_ca_topic_score_codex":0.011657179,"about_ca_topic_score_gemma":0.008183736,"teacher_disagreement_score":0.011657179,"about_ca_system_score_codex":0.0019037259,"about_ca_system_score_gemma":0.002182624,"threshold_uncertainty_score":0.023178637},"labels":[],"label_agreement":null},{"id":"W2786472750","doi":"10.1609/aaai.v32i1.11332","title":"Complex Sequential Question Answering: Towards Learning to Converse Over Linked Question Answer Pairs with a Knowledge Graph","year":2018,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":177,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Question answering; Converse; Dialog box; Ask price; Conversation; Natural language processing; Context (archaeology); Artificial intelligence; Coreference; Ellipsis (linguistics); Parsing; Graph; Information retrieval; Linguistics; World Wide Web; Theoretical computer science; Resolution (logic); Mathematics","score_opus":0.09260613830365648,"score_gpt":0.32629793119242484,"score_spread":0.23369179288876835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786472750","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07319394,0.0014016053,0.9012562,0.0019496941,0.000118539756,0.0006712001,0.0047854693,0.012062462,0.0045607793],"genre_scores_gemma":[0.41988757,0.00045370474,0.55112576,0.0016534242,0.0001511853,0.0005718339,0.02017256,0.00046583233,0.0055181594],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969944,0.0010658769,0.00012052752,0.001423782,0.0002760671,0.00011939094],"domain_scores_gemma":[0.99164987,0.0059987027,0.00041037443,0.0011295193,0.00050581637,0.00030577485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028726726,0.0022935835,0.00097450975,0.0021021252,0.0009145248,0.0018554695,0.0038558508,0.0032023126,0.0055846423],"category_scores_gemma":[0.01301584,0.0007073598,0.0023120125,0.0012701004,0.0012385531,0.006546896,0.0036387448,0.0035823614,0.0032542048],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010535553,0.0014382366,0.01269185,0.0018855047,0.00057441613,0.001212337,0.004485568,0.19447358,0.020458579,0.032264743,0.047119573,0.68234205],"study_design_scores_gemma":[0.00010958571,0.00027962404,0.0019697032,0.00011255751,0.00012238724,0.0003296854,0.0009529556,0.8673593,0.006260517,0.10496183,0.017485628,0.00005628544],"about_ca_topic_score_codex":0.0055745197,"about_ca_topic_score_gemma":0.010439543,"teacher_disagreement_score":0.0055846423,"about_ca_system_score_codex":0.0014861256,"about_ca_system_score_gemma":0.0014023788,"threshold_uncertainty_score":0.01868242},"labels":[],"label_agreement":null},{"id":"W2786983967","doi":"10.24963/ijcai.2018/609","title":"An Ensemble of Retrieval-Based and Generation-Based Human-Computer Conversation Systems","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":114,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Natural Science Foundation of China","keywords":"Computer science; Conversation; Ranking (information retrieval); Natural language generation; Generative grammar; Margin (machine learning); Utterance; Generator (circuit theory); Artificial intelligence; Process (computing); Information retrieval; Natural language processing; Artificial neural network; Natural language; Machine learning; Programming language","score_opus":0.03805082076957427,"score_gpt":0.27124217934459155,"score_spread":0.23319135857501727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786983967","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09013939,0.0020250073,0.89312977,0.0006658401,0.00029913313,0.000364149,0.00022790945,0.0076170024,0.0055317604],"genre_scores_gemma":[0.72307,0.00066952256,0.26636615,0.00038373706,0.00035858466,0.00028060778,0.00081573526,0.00025617043,0.0077993865],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99855596,0.0003883099,0.00008571724,0.00052147574,0.00030923303,0.00013926384],"domain_scores_gemma":[0.9983498,0.0005055644,0.00009628451,0.00030644154,0.0005602127,0.00018176313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028053962,0.0012708592,0.0014627422,0.0013379208,0.0011261619,0.0012264217,0.001586033,0.0013384178,0.002017026],"category_scores_gemma":[0.004415407,0.0006204631,0.0011113235,0.0009298092,0.00046883372,0.0028786957,0.002197324,0.0012674001,0.0014946922],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007526788,0.00058948627,0.010593827,0.000342388,0.0005862039,0.0005141799,0.0011661265,0.10718052,0.053893443,0.0056386827,0.008327752,0.81041473],"study_design_scores_gemma":[0.000022739552,0.00020818376,0.0020025813,0.000018475035,0.00019461628,0.00027843352,0.00010549898,0.97957647,0.008769693,0.0036737772,0.005099832,0.000049651033],"about_ca_topic_score_codex":0.004883686,"about_ca_topic_score_gemma":0.005115542,"teacher_disagreement_score":0.004883686,"about_ca_system_score_codex":0.0005832874,"about_ca_system_score_gemma":0.0011599051,"threshold_uncertainty_score":0.01483655},"labels":[],"label_agreement":null},{"id":"W2787436795","doi":"10.1037/rev0000133","title":"A model of event knowledge.","year":2019,"lang":"en","type":"article","venue":"Psychological Review","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Connectionism; Computer science; Event (particle physics); Cognitive science; Terminology; Complex event processing; Knowledge representation and reasoning; Representation (politics); Artificial intelligence; Scripting language; Data science; Natural language processing; Psychology; Linguistics; Artificial neural network","score_opus":0.1459229846289665,"score_gpt":0.4090688148763844,"score_spread":0.2631458302474179,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2787436795","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010177605,0.007994335,0.8278689,0.010351972,0.00042227522,0.0002759552,0.0025307632,0.00090917957,0.13946895],"genre_scores_gemma":[0.57936877,0.016493158,0.34564894,0.0016958816,0.0010632817,0.0011495095,0.005626447,0.00030216345,0.04865189],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9988625,0.00034994533,0.00010258746,0.0004149478,0.00020756265,0.00006241756],"domain_scores_gemma":[0.99713206,0.0018454007,0.00034195877,0.0002743053,0.00027127922,0.00013501168],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017018142,0.00089248543,0.0005511743,0.00307532,0.0009851472,0.004294612,0.0026073188,0.0028060954,0.016411312],"category_scores_gemma":[0.0062105763,0.00067041215,0.0019208933,0.0032249773,0.0032234786,0.0118542155,0.0016215652,0.0021696268,0.0035885451],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017470962,0.000014652908,0.00057011185,0.00018329635,0.000040585423,0.00010845738,0.00046280134,0.012591756,0.00014438474,0.96363604,0.0032405853,0.018989759],"study_design_scores_gemma":[0.00002359696,0.000013320006,0.0005198661,0.00016650048,0.000047259968,0.00033198166,0.00014856047,0.055117395,0.00016330872,0.8742588,0.06919396,0.00001541308],"about_ca_topic_score_codex":0.008455156,"about_ca_topic_score_gemma":0.0033545163,"teacher_disagreement_score":0.016411312,"about_ca_system_score_codex":0.0034888145,"about_ca_system_score_gemma":0.0024634781,"threshold_uncertainty_score":0.0549013},"labels":[],"label_agreement":null},{"id":"W2787882523","doi":"10.1145/3178876.3186017","title":"Scalable Instance Reconstruction in Knowledge Bases via Relatedness Affiliated Embedding","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Scalability; Embedding; Computer science; Computational complexity theory; Time complexity; Reduction (mathematics); Schema (genetic algorithms); Theoretical computer science; Computation; Algorithm; Artificial intelligence; Mathematics; Machine learning","score_opus":0.02146768690508039,"score_gpt":0.2696430564403868,"score_spread":0.2481753695353064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2787882523","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034269243,0.0001986725,0.9624435,0.0002658481,0.000013056432,0.00006107682,0.00031145863,0.0018106882,0.00062643416],"genre_scores_gemma":[0.3624342,0.00033248856,0.6308156,0.00018365124,0.00005002178,0.00015287932,0.0033250898,0.0003285444,0.002377501],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99838924,0.000526165,0.00009145338,0.0004995608,0.00035988173,0.00013370195],"domain_scores_gemma":[0.99494535,0.002814548,0.00036452874,0.0013283115,0.000401953,0.00014533885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019873022,0.00077520113,0.001359845,0.0013456187,0.00055249996,0.0020447953,0.0020377608,0.0015005714,0.0021575268],"category_scores_gemma":[0.010377456,0.00076684466,0.0013418953,0.0020376726,0.0009477753,0.005315792,0.0027019188,0.0024807304,0.0009982465],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002984054,0.00022446945,0.0028926316,0.00026307045,0.00012856415,0.00023009488,0.00045045064,0.5592651,0.0067628943,0.030919522,0.006045239,0.3925195],"study_design_scores_gemma":[0.000010664857,0.00002102922,0.00020089324,0.000008318784,0.000013968565,0.000043189957,0.00006052965,0.97703105,0.0021928765,0.01958026,0.0008294805,0.000007775582],"about_ca_topic_score_codex":0.0054200552,"about_ca_topic_score_gemma":0.0055732164,"teacher_disagreement_score":0.0054200552,"about_ca_system_score_codex":0.0009839695,"about_ca_system_score_gemma":0.0010442059,"threshold_uncertainty_score":0.010776997},"labels":[],"label_agreement":null},{"id":"W2788343755","doi":"10.1609/aaai.v32i1.11273","title":"CA-RNN: Using Context-Aligned Recurrent Neural Networks for Modeling Sentence Similarity","year":2018,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Recurrent neural network; Computer science; Sentence; Benchmark (surveying); Context (archaeology); Artificial intelligence; Paraphrase; Natural language processing; Word (group theory); Similarity (geometry); Language model; Artificial neural network; Speech recognition","score_opus":0.17887027404759642,"score_gpt":0.3348507905021531,"score_spread":0.15598051645455668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2788343755","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07135659,0.0025950347,0.91563857,0.00038103675,0.00033609226,0.00020891661,0.0009685861,0.004853109,0.0036620367],"genre_scores_gemma":[0.7229157,0.0010131766,0.26625922,0.0004554415,0.00020860264,0.00038875104,0.0023885593,0.0003848445,0.0059857094],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99952745,0.00011790502,0.000032620057,0.00019726816,0.00008033802,0.000044410303],"domain_scores_gemma":[0.999371,0.00024320043,0.00008353847,0.000067297464,0.00020657777,0.000028439992],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00087246235,0.0010077327,0.0007083397,0.00085330143,0.00034453676,0.00061598007,0.0019739922,0.0009441165,0.0017689073],"category_scores_gemma":[0.003226305,0.00037813018,0.0008209418,0.00084427017,0.0003280387,0.0014891494,0.00065026194,0.00119536,0.00078212144],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043529668,0.00021916385,0.003899921,0.00026321466,0.000345704,0.00034403274,0.00023879147,0.5238455,0.022938881,0.009333345,0.008954878,0.42918134],"study_design_scores_gemma":[0.0000056886965,0.000023568966,0.00022582004,0.0000057550783,0.000017483579,0.000018201417,0.000004940762,0.9964006,0.001009637,0.0017692959,0.0005127093,0.000006371658],"about_ca_topic_score_codex":0.018162297,"about_ca_topic_score_gemma":0.022042392,"teacher_disagreement_score":0.018162297,"about_ca_system_score_codex":0.0008752813,"about_ca_system_score_gemma":0.0010191551,"threshold_uncertainty_score":0.036113203},"labels":[],"label_agreement":null},{"id":"W2788463229","doi":"10.18653/v1/n18-1128","title":"Reusing Weights in Subword-Aware Neural Language Models","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Ministry of Education and Science of the Republic of Kazakhstan; Nvidia","keywords":"Reuse; Syllable; Morpheme; Computer science; Margin (machine learning); Embedding; Word (group theory); Layer (electronics); Aggregate (composite); Character (mathematics); Simple (philosophy); Artificial intelligence; Language model; Natural language processing; Speech recognition; Machine learning; Linguistics; Mathematics; Engineering","score_opus":0.04677271909570187,"score_gpt":0.2847590447914972,"score_spread":0.23798632569579536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2788463229","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0708045,0.0032474601,0.91946554,0.00071380136,0.00036217013,0.00006406513,0.00030511513,0.0026239373,0.0024133166],"genre_scores_gemma":[0.8157458,0.0019272263,0.17298701,0.0003298281,0.00040596738,0.00018823559,0.0011450045,0.00072505814,0.0065459046],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995165,0.00015787724,0.000044726206,0.00015066603,0.00006999082,0.000060332804],"domain_scores_gemma":[0.9980046,0.0012308717,0.00009330046,0.00025654607,0.00033872447,0.000076044824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011019586,0.0013361559,0.0013122557,0.0010339209,0.00041159004,0.0015907647,0.0021569715,0.0015066866,0.0023399666],"category_scores_gemma":[0.007059255,0.00095553574,0.0008061531,0.001305123,0.00046947502,0.00443338,0.0016731485,0.00263237,0.0016815028],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044638623,0.0002339192,0.0012281494,0.00023675141,0.0003281017,0.00011552364,0.0002035161,0.46415156,0.0088863615,0.011379674,0.0055036326,0.5072864],"study_design_scores_gemma":[0.000008044327,0.00001563135,0.000058644437,0.0000069455045,0.000030706673,0.000011346516,0.000008072016,0.9885484,0.0008695086,0.010154083,0.0002828706,0.0000057975],"about_ca_topic_score_codex":0.0051977444,"about_ca_topic_score_gemma":0.008413316,"teacher_disagreement_score":0.0051977444,"about_ca_system_score_codex":0.00057923346,"about_ca_system_score_gemma":0.0008280764,"threshold_uncertainty_score":0.010334969},"labels":[],"label_agreement":null},{"id":"W2788918109","doi":"10.1609/aaai.v32i1.12015","title":"Spectral Word Embedding with Negative Sampling","year":2018,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Word (group theory); Embedding; Computer science; Word embedding; Context (archaeology); Sampling (signal processing); Algorithm; Artificial intelligence; Natural language processing; Theoretical computer science; Mathematics","score_opus":0.10624167690963107,"score_gpt":0.32364312890205293,"score_spread":0.21740145199242186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2788918109","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020199519,0.0002854063,0.97783285,0.0001822306,0.00005380971,0.00004367166,0.000042750053,0.0004771886,0.0008825477],"genre_scores_gemma":[0.4648861,0.0003146439,0.52858156,0.00041093657,0.0002324805,0.00025575343,0.00060912006,0.0003018391,0.0044075707],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982293,0.0009885269,0.000069307986,0.0003311289,0.00030291948,0.000078776044],"domain_scores_gemma":[0.99505144,0.0032087136,0.00026481962,0.0008145219,0.0005465761,0.000113833004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022115167,0.0010676013,0.0011676931,0.0011120087,0.0005410327,0.0010218273,0.0013842375,0.0012595268,0.002285907],"category_scores_gemma":[0.011906762,0.00047907844,0.00062777055,0.0009906051,0.0014819398,0.0039014327,0.0017366755,0.0013801857,0.0011374311],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041629997,0.0003657512,0.0027580392,0.00039217752,0.00012594432,0.00016409467,0.00033604406,0.3772026,0.011435448,0.10174599,0.0064497143,0.49860787],"study_design_scores_gemma":[0.000013706486,0.000030907962,0.00011506625,0.00000880923,0.000005613491,0.000044993685,0.000020098296,0.96174663,0.001505516,0.03558368,0.00091723196,0.000007775097],"about_ca_topic_score_codex":0.00084210455,"about_ca_topic_score_gemma":0.0012538891,"teacher_disagreement_score":0.002285907,"about_ca_system_score_codex":0.000484163,"about_ca_system_score_gemma":0.00049585936,"threshold_uncertainty_score":0.011695743},"labels":[],"label_agreement":null},{"id":"W2791007156","doi":"10.1111/cogs.12583","title":"A Large‐Scale Analysis of Variance in Written Language","year":2018,"lang":"en","type":"article","venue":"Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Natural language processing; Semantics (computer science); Artificial intelligence; Natural language; Distributional semantics; Universal Networking Language; Language identification; Language model; Linguistics; Word (group theory); Comprehension approach; Semantic similarity; Programming language","score_opus":0.017755877792834312,"score_gpt":0.3001291638628089,"score_spread":0.2823732860699746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2791007156","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.71324587,0.0020029766,0.26810044,0.001425299,0.0003422122,0.00022239785,0.0025916204,0.0011230058,0.010946175],"genre_scores_gemma":[0.96891886,0.00029899718,0.02759529,0.00013923332,0.00015928067,0.00016403175,0.0013397634,0.0001384187,0.001246088],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99661356,0.0017049346,0.000121124795,0.0008116187,0.0006752604,0.00007337921],"domain_scores_gemma":[0.97911227,0.016032489,0.0010406637,0.002822065,0.0008042447,0.00018835353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00401518,0.00036748862,0.00045047037,0.0021158948,0.00054330885,0.0014917871,0.00042319475,0.00041195055,0.002161428],"category_scores_gemma":[0.030860033,0.00017148863,0.0008275526,0.0025967874,0.0010314254,0.0013242364,0.0010590451,0.0011654227,0.00051248376],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085992616,0.000650473,0.29542282,0.0006295508,0.0027001433,0.0008585499,0.009242698,0.013493298,0.032775637,0.04891321,0.021850577,0.5726031],"study_design_scores_gemma":[0.000057990266,0.00051306584,0.8293188,0.00011973343,0.00035148146,0.0009171937,0.0019100779,0.072190195,0.0054258094,0.06584558,0.023189478,0.00016056382],"about_ca_topic_score_codex":0.0017396798,"about_ca_topic_score_gemma":0.0018189473,"teacher_disagreement_score":0.00401518,"about_ca_system_score_codex":0.00042135673,"about_ca_system_score_gemma":0.0003540826,"threshold_uncertainty_score":0.021234512},"labels":[],"label_agreement":null},{"id":"W2791167821","doi":"10.1162/neco_a_01077","title":"Facet Annotation by Extending CNN with a Matching Strategy","year":2018,"lang":"en","type":"article","venue":"Neural Computation","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Facet (psychology); Matching (statistics); Mathematics; Representation (politics); Convolutional neural network; Pattern recognition (psychology); Computer science; Similarity (geometry); Artificial intelligence; Image (mathematics); Statistics; Psychology","score_opus":0.027028736086365145,"score_gpt":0.28309684037867144,"score_spread":0.2560681042923063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2791167821","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09603797,0.0007347085,0.8886673,0.00035249782,0.00011367581,0.00018187461,0.0010073514,0.0064941356,0.0064104875],"genre_scores_gemma":[0.73903406,0.00041283655,0.24703157,0.00042612423,0.00010248274,0.00022107456,0.003300841,0.00028566233,0.009185308],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99971515,0.000034034776,0.000017890601,0.00013110587,0.00005562474,0.00004612161],"domain_scores_gemma":[0.9996402,0.00008516114,0.000037329388,0.00010623176,0.00010560095,0.000025498806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00047918988,0.0009315064,0.0004977071,0.00094899104,0.00038473617,0.0006429101,0.0010482866,0.00084175594,0.0020610609],"category_scores_gemma":[0.0016729509,0.00028937572,0.0008551585,0.0012024449,0.00043148178,0.0019486961,0.001199047,0.00064424414,0.0010287013],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003645091,0.00022902818,0.0075981487,0.00022255107,0.00013139132,0.00024709213,0.0004729296,0.09702206,0.038540564,0.009862797,0.014422334,0.8308866],"study_design_scores_gemma":[0.000009485221,0.000039472263,0.001292501,0.000010990611,0.00003740389,0.000065883265,0.000043852026,0.9786789,0.008837966,0.0073810397,0.0035897673,0.000012699361],"about_ca_topic_score_codex":0.017096981,"about_ca_topic_score_gemma":0.01979385,"teacher_disagreement_score":0.017096981,"about_ca_system_score_codex":0.0008830955,"about_ca_system_score_gemma":0.00061655947,"threshold_uncertainty_score":0.033994973},"labels":[],"label_agreement":null},{"id":"W2796265348","doi":"10.1007/978-3-319-89656-4_31","title":"Matching Résumés to Job Descriptions with Stacked Models","year":2018,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Matching (statistics); Task (project management); Artificial intelligence; Artificial neural network; Data mining; Machine learning; Operations research; Information retrieval; Statistics; Management; Mathematics","score_opus":0.04430563088816648,"score_gpt":0.2630630116354944,"score_spread":0.21875738074732792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2796265348","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02433424,0.00055002695,0.91742545,0.00084724417,0.0009978035,0.00052103575,0.0058883107,0.021255065,0.02818077],"genre_scores_gemma":[0.36633536,0.00095087365,0.52702206,0.000550692,0.00042954122,0.00050889753,0.018580068,0.006657266,0.07896527],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99769443,0.0004875196,0.00018547637,0.0005269424,0.00078577927,0.0003199031],"domain_scores_gemma":[0.99495584,0.0017390836,0.00023422849,0.001957406,0.00077295146,0.00034040285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028852113,0.0011396958,0.0012065121,0.0020850473,0.0010766581,0.004443582,0.0029313036,0.0016241211,0.03570674],"category_scores_gemma":[0.018001819,0.0011350217,0.0024651415,0.0023229118,0.00084645563,0.005813336,0.0033272963,0.0020801637,0.01686398],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016585744,0.00077691756,0.0049950043,0.00075681525,0.00017503515,0.00065200165,0.0012983807,0.13481666,0.008011451,0.20908141,0.09782947,0.5399483],"study_design_scores_gemma":[0.00006337517,0.00012967864,0.0010642378,0.00016036067,0.00010500995,0.00014153899,0.00054283964,0.60658145,0.01238989,0.2913123,0.08741514,0.00009424403],"about_ca_topic_score_codex":0.0063167913,"about_ca_topic_score_gemma":0.006775576,"teacher_disagreement_score":0.03570674,"about_ca_system_score_codex":0.0015089278,"about_ca_system_score_gemma":0.0021465793,"threshold_uncertainty_score":0.11945093},"labels":[],"label_agreement":null},{"id":"W2798416089","doi":"10.18653/v1/p18-1224","title":"Neural Natural Language Inference Models Enhanced with External Knowledge","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":276,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Queen's University","funders":"","keywords":"Computer science; Inference; Leverage (statistics); Artificial intelligence; Artificial neural network; Natural language; Machine learning; Task (project management); Language model; Natural language processing","score_opus":0.02114323554206217,"score_gpt":0.2839945996636008,"score_spread":0.26285136412153864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2798416089","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04581743,0.0027908485,0.93347305,0.0020310404,0.00021712009,0.00014101446,0.0032496427,0.0055280444,0.0067518447],"genre_scores_gemma":[0.65426326,0.0017345169,0.31222838,0.0009993066,0.00050935626,0.00056820846,0.01787026,0.0005600259,0.011266725],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99884844,0.00038058154,0.00006970411,0.00046617462,0.0001464876,0.00008865963],"domain_scores_gemma":[0.9943321,0.003983988,0.0002861684,0.00071112305,0.000567384,0.00011922723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030584824,0.0012366698,0.0010930763,0.001244867,0.0006000409,0.0016547753,0.002749116,0.0014063697,0.003623985],"category_scores_gemma":[0.011683591,0.0005331033,0.0014498865,0.0013834026,0.0007375263,0.0056374697,0.0015952566,0.004252959,0.0018968162],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029850978,0.00041130363,0.004725055,0.00040095721,0.00032180603,0.00027058407,0.00036470042,0.63483024,0.0031436114,0.03346653,0.020207068,0.30155963],"study_design_scores_gemma":[0.000007623727,0.000009348512,0.0002101642,0.000011038838,0.000020400654,0.00001770875,0.000009387268,0.9872328,0.00045594026,0.010821035,0.0011984186,0.000006094028],"about_ca_topic_score_codex":0.01136504,"about_ca_topic_score_gemma":0.020949107,"teacher_disagreement_score":0.01136504,"about_ca_system_score_codex":0.0013169266,"about_ca_system_score_gemma":0.0011754065,"threshold_uncertainty_score":0.02259773},"labels":[],"label_agreement":null},{"id":"W2798934669","doi":"10.1145/3209978.3210140","title":"What Do Viewers Say to Their TVs?","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Domain (mathematical analysis); Taxonomy (biology); Information retrieval; Entertainment; Service (business); World Wide Web; Product (mathematics)","score_opus":0.028074258477853258,"score_gpt":0.26552990712584185,"score_spread":0.2374556486479886,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2798934669","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8769061,0.0030125624,0.049573097,0.0048991414,0.00024384599,0.00016952396,0.01737843,0.0014718844,0.046345454],"genre_scores_gemma":[0.9850696,0.00072241144,0.005073549,0.0003622,0.00015340594,0.000037234866,0.0039942805,0.00011654405,0.0044709058],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9995524,0.0001587562,0.000028157712,0.00010313168,0.00008770512,0.00006979969],"domain_scores_gemma":[0.99710184,0.0019855793,0.0003069339,0.00013101005,0.0003361071,0.00013848588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006999146,0.00029021758,0.00025518713,0.0009319289,0.00039697328,0.0014980179,0.00023707743,0.00047924946,0.0034112127],"category_scores_gemma":[0.005332145,0.000106773565,0.00026528668,0.0009833706,0.00034415707,0.0018101336,0.000389979,0.00053004985,0.0011215108],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016121188,0.00031361202,0.5107556,0.0011088724,0.00019257667,0.00097358617,0.040752623,0.0049076034,0.032933235,0.039475463,0.07603887,0.29093587],"study_design_scores_gemma":[0.00007788621,0.0005762242,0.5285301,0.000431161,0.00040899144,0.0030065435,0.05967087,0.075369455,0.02465437,0.039810672,0.26722014,0.00024351131],"about_ca_topic_score_codex":0.010116559,"about_ca_topic_score_gemma":0.010195875,"teacher_disagreement_score":0.010116559,"about_ca_system_score_codex":0.0004677278,"about_ca_system_score_gemma":0.0002841111,"threshold_uncertainty_score":0.020115316},"labels":[],"label_agreement":null},{"id":"W2799060650","doi":"10.18653/v1/p18-2089","title":"Addressing Noise in Multidialectal Word Embeddings","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; New York University Abu Dhabi","keywords":"Computer science; Natural language processing; Artificial intelligence; Boosting (machine learning); Sentence; Word (group theory); Embedding; Speech recognition; Spelling; Task (project management); Noise (video); Linguistics","score_opus":0.05947592401808097,"score_gpt":0.3114900301775979,"score_spread":0.2520141061595169,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2799060650","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17978168,0.0013473697,0.8041962,0.00075396046,0.0005217202,0.00018238866,0.000962969,0.009603278,0.00265042],"genre_scores_gemma":[0.628977,0.00075176475,0.35557538,0.0005649633,0.000289551,0.00026144166,0.0060226624,0.001772325,0.0057848166],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99568117,0.001870815,0.00036861884,0.0011715821,0.00072436396,0.00018339188],"domain_scores_gemma":[0.98546195,0.007660795,0.000664506,0.004142582,0.0017952122,0.00027495233],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003611872,0.0020971897,0.0014405991,0.0014914306,0.0008395759,0.00229795,0.0013701421,0.0016019882,0.002382676],"category_scores_gemma":[0.026976602,0.0005876494,0.0009307643,0.0018288218,0.0008574504,0.005993343,0.0042539993,0.0026648676,0.004083037],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008959395,0.00058475323,0.01384128,0.00083793345,0.00030583542,0.0004309714,0.0016139535,0.045685265,0.06224244,0.010643463,0.01787312,0.84504503],"study_design_scores_gemma":[0.000090538815,0.000590908,0.005888468,0.00017315605,0.0001858785,0.0013120079,0.001216506,0.8349933,0.08490719,0.04929381,0.021219103,0.00012918425],"about_ca_topic_score_codex":0.0010528368,"about_ca_topic_score_gemma":0.0023678208,"teacher_disagreement_score":0.003611872,"about_ca_system_score_codex":0.00045684865,"about_ca_system_score_gemma":0.00080438936,"threshold_uncertainty_score":0.01910168},"labels":[],"label_agreement":null},{"id":"W2800110000","doi":"10.1016/j.ipm.2018.04.007","title":"Neural word and entity embeddings for ad hoc retrieval","year":2018,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Task (project management); Word (group theory); Information retrieval; Artificial intelligence; Natural language processing; Question answering; Artificial neural network; Mathematics","score_opus":0.016156215218341936,"score_gpt":0.26573830482327654,"score_spread":0.2495820896049346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2800110000","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09075751,0.005960549,0.88844126,0.0015731724,0.00069933373,0.00020951078,0.0034356557,0.0052556433,0.0036674142],"genre_scores_gemma":[0.66234356,0.0031842142,0.30385312,0.0004933644,0.0007233856,0.0003094789,0.012700032,0.00044118115,0.015951626],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991248,0.00026763525,0.0000965226,0.00024222219,0.00014881603,0.00012012134],"domain_scores_gemma":[0.9980463,0.0009174833,0.00013223122,0.0004439453,0.00035672003,0.00010343247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015429696,0.0010823826,0.0012947257,0.0027652609,0.0005990386,0.0016760567,0.0016577756,0.0018419281,0.0037972918],"category_scores_gemma":[0.005424967,0.0005533234,0.0010408022,0.0032083215,0.0005489298,0.0053829006,0.0019394757,0.00183918,0.0026003243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077746273,0.00058159337,0.0030929416,0.0003858845,0.0003021808,0.00021185257,0.00026928546,0.07560748,0.010008844,0.016926145,0.028149404,0.863687],"study_design_scores_gemma":[0.000043797547,0.000121591416,0.0009293232,0.000035267207,0.0000921343,0.00012338908,0.0001303104,0.9461066,0.0033577592,0.0447432,0.004285979,0.000030754738],"about_ca_topic_score_codex":0.0054656314,"about_ca_topic_score_gemma":0.007942118,"teacher_disagreement_score":0.0054656314,"about_ca_system_score_codex":0.0009112196,"about_ca_system_score_gemma":0.001262289,"threshold_uncertainty_score":0.01270318},"labels":[],"label_agreement":null},{"id":"W28026516","doi":"","title":"Sparse Matrix Factorization: Applications to Latent Semantic Indexing","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Singular value decomposition; Matrix decomposition; Computer science; Ranking (information retrieval); Sparse matrix; Latent semantic analysis; Selection (genetic algorithm); Search engine indexing; Column (typography); Factorization; Information retrieval; Task (project management); Matrix (chemical analysis); Data mining; Artificial intelligence; Algorithm; Engineering","score_opus":0.03445608389451069,"score_gpt":0.28726715167066513,"score_spread":0.2528110677761544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W28026516","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0028231791,0.0024475234,0.9908568,0.00062552025,0.00010280488,0.00009063989,0.00031219464,0.0011006861,0.001640674],"genre_scores_gemma":[0.09564229,0.004724508,0.89440304,0.00021224338,0.00076925155,0.00027860366,0.0013365211,0.00017700481,0.0024565628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998387,0.0006489862,0.00011615558,0.00024399845,0.00051866163,0.00008516557],"domain_scores_gemma":[0.9954905,0.0032003922,0.00027691937,0.0005057895,0.00043068037,0.00009586683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00360708,0.0011089841,0.0013311425,0.0038527274,0.0009573232,0.0021437092,0.0010340716,0.0010943599,0.0030792535],"category_scores_gemma":[0.010633833,0.0005610759,0.0013654785,0.0063849953,0.0010371277,0.002467395,0.0015958027,0.0018996557,0.001581116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013769123,0.00015670729,0.0020284757,0.00055495836,0.00017575223,0.00019137438,0.00058498984,0.06682126,0.007672776,0.10691562,0.016880482,0.7978799],"study_design_scores_gemma":[0.000052448242,0.000067418834,0.0009805744,0.00007983411,0.00005276335,0.00032571398,0.00014618767,0.7456113,0.0046890154,0.22323719,0.024684276,0.00007330483],"about_ca_topic_score_codex":0.0058934595,"about_ca_topic_score_gemma":0.0070855613,"teacher_disagreement_score":0.0058934595,"about_ca_system_score_codex":0.00092335383,"about_ca_system_score_gemma":0.0010792855,"threshold_uncertainty_score":0.019076288},"labels":[],"label_agreement":null},{"id":"W2803395921","doi":"10.18653/v1/s18-2001","title":"Resolving Event Coreference with Supervised Representation Learning and Clustering-Oriented Regularization","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Coreference; Cluster analysis; Computer science; Artificial intelligence; Regularization (linguistics); Event (particle physics); Representation (politics); Machine learning; Natural language processing; Feature learning; Resolution (logic)","score_opus":0.030392454071176973,"score_gpt":0.2756032686607547,"score_spread":0.24521081458957772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2803395921","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0052236966,0.000093197705,0.99347645,0.0001909654,0.000018152588,0.000024154968,0.000036046167,0.00029997152,0.0006374145],"genre_scores_gemma":[0.30947578,0.00026047335,0.68327874,0.0004908625,0.00019316119,0.00025138923,0.00086754916,0.00047671545,0.0047053355],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99593973,0.0019124774,0.00017959098,0.0011762088,0.0005533607,0.00023859105],"domain_scores_gemma":[0.9928046,0.0036221163,0.0007475854,0.001991316,0.00067921606,0.00015517481],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00538517,0.0013274065,0.001500745,0.0020824138,0.0016259071,0.0022225154,0.0041041113,0.0034917716,0.0022695356],"category_scores_gemma":[0.017290644,0.0010174254,0.0016293188,0.0030505452,0.002131222,0.005419919,0.0050485847,0.004758193,0.0009889572],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002442392,0.00027937372,0.002486768,0.00028324383,0.00033383112,0.00033147592,0.0013806237,0.47497368,0.010778742,0.17460822,0.011257195,0.3230426],"study_design_scores_gemma":[0.000009974639,0.000017215018,0.00016138004,0.000012543628,0.000016173894,0.000049267626,0.000056774923,0.91352075,0.0022571941,0.08222176,0.0016601095,0.00001686327],"about_ca_topic_score_codex":0.0031866112,"about_ca_topic_score_gemma":0.00488698,"teacher_disagreement_score":0.00538517,"about_ca_system_score_codex":0.0016724253,"about_ca_system_score_gemma":0.0016767388,"threshold_uncertainty_score":0.028479874},"labels":[],"label_agreement":null},{"id":"W2803521238","doi":"10.1007/978-3-319-91464-0_30","title":"Deep Learning in Automated Essay Scoring","year":2018,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Hyperparameter; Computer science; Artificial intelligence; Suite; Machine learning; Metric (unit); Deep learning; Exploit; Convolutional neural network; Kappa; Set (abstract data type); Artificial neural network; Construct (python library); Feature engineering; Sample (material); Mathematics; Programming language","score_opus":0.01841032280520948,"score_gpt":0.2519021637832037,"score_spread":0.2334918409779942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2803521238","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03298436,0.009753668,0.93708336,0.0019557723,0.00075674726,0.000073486946,0.0005350357,0.0029145437,0.013943025],"genre_scores_gemma":[0.6287863,0.004063049,0.31728235,0.00039336033,0.00096927397,0.00022320381,0.0018476164,0.0005696687,0.045865234],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99812347,0.0010320256,0.00010251292,0.00020131595,0.00038014722,0.00016050535],"domain_scores_gemma":[0.99474585,0.003340906,0.00022647572,0.0005335741,0.0009938003,0.00015930683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00298766,0.0008073201,0.00091817917,0.0013111775,0.00045397785,0.0016287534,0.0013851296,0.0011004851,0.0053213695],"category_scores_gemma":[0.010670084,0.0005484542,0.0004260474,0.0015413851,0.0005053048,0.002613361,0.0021084424,0.002476068,0.0031108356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008077683,0.00010647919,0.0018196424,0.00012358518,0.00004432723,0.00002807016,0.00008493569,0.04007541,0.00123292,0.013932143,0.01965544,0.9228164],"study_design_scores_gemma":[0.000011670641,0.000040400388,0.0012667297,0.00009922352,0.000022388522,0.00004552234,0.000048684837,0.921741,0.0027026748,0.06422117,0.0097800335,0.000020537751],"about_ca_topic_score_codex":0.0023432728,"about_ca_topic_score_gemma":0.0051275417,"teacher_disagreement_score":0.0053213695,"about_ca_system_score_codex":0.00082318956,"about_ca_system_score_gemma":0.0009936761,"threshold_uncertainty_score":0.017801821},"labels":[],"label_agreement":null},{"id":"W2804168033","doi":"10.1002/asi.24046","title":"In‐text function of author self‐citations: Implications for research evaluation practice","year":2018,"lang":"en","type":"article","venue":"Journal of the Association for Information Science and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Universiteit Leiden","keywords":"Citation; Computer science; Information retrieval; Function (biology); Coding (social sciences); Citation analysis; Field (mathematics); Data science; Psychology; Statistics; Library science; Mathematics","score_opus":0.09869317634340848,"score_gpt":0.4410832076101872,"score_spread":0.34239003126677875,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2804168033","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5631347,0.034237932,0.29899082,0.033778634,0.0032991434,0.008353244,0.0056545427,0.0017432964,0.050807793],"genre_scores_gemma":[0.9408082,0.0012003828,0.051521286,0.0009587422,0.00036847772,0.0028729453,0.0008918607,0.0001975558,0.0011806174],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.42411327,0.47816995,0.042560004,0.013366876,0.039258882,0.0025311098],"domain_scores_gemma":[0.061564125,0.8517317,0.025097508,0.018935986,0.04030565,0.0023650106],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.56287926,0.001691622,0.004001479,0.019785076,0.0026147983,0.015818592,0.00475715,0.0034723897,0.0051824334],"category_scores_gemma":[0.87906474,0.00062624435,0.0030091773,0.023607196,0.0060118046,0.023586877,0.0052927113,0.0031286308,0.0011830995],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029998245,0.0009303251,0.5567661,0.00947137,0.005563631,0.00026110793,0.018301021,0.005844439,0.000675843,0.044941533,0.011965143,0.34227967],"study_design_scores_gemma":[0.0012235767,0.0042004874,0.4666535,0.01223582,0.0058390065,0.0009997868,0.030487506,0.15201025,0.005817352,0.2876026,0.03233997,0.0005902788],"about_ca_topic_score_codex":0.0040394356,"about_ca_topic_score_gemma":0.003253354,"teacher_disagreement_score":0.98021495,"about_ca_system_score_codex":0.0077377423,"about_ca_system_score_gemma":0.008129913,"threshold_uncertainty_score":0.5390477},"labels":[],"label_agreement":null},{"id":"W2804709417","doi":"10.5339/qfarc.2018.ictpd881","title":"Towards OpenDomain CrossLanguage Question Answering","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Question answering; Computer science; Natural language processing; Task (project management); Machine translation; Open domain; Artificial intelligence; Domain (mathematical analysis); World Wide Web; Information retrieval; Semitic languages; Natural language; Arabic; Linguistics","score_opus":0.018347693361357918,"score_gpt":0.2942623266627611,"score_spread":0.27591463330140314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2804709417","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023924029,0.0017985568,0.9347531,0.0033263816,0.00036385815,0.0003800592,0.0041460013,0.02114621,0.010161725],"genre_scores_gemma":[0.17866237,0.0007482955,0.77702504,0.0031001584,0.00036292273,0.0006624127,0.031421483,0.0013933973,0.006623909],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99052936,0.004483815,0.0007889727,0.0022396392,0.0014059119,0.00055232923],"domain_scores_gemma":[0.98597336,0.0072235023,0.00056234136,0.0027926953,0.002832187,0.00061592774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007833936,0.0015286236,0.0012533106,0.0030455652,0.0016193149,0.0047905133,0.002566626,0.003396643,0.008759454],"category_scores_gemma":[0.021505047,0.0007309267,0.002176666,0.0023043796,0.0019139952,0.01116421,0.011811373,0.0045034345,0.006852683],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008487024,0.0013352993,0.007896926,0.0030251013,0.00042645168,0.0022976345,0.008603522,0.023686463,0.03832153,0.18050678,0.10955675,0.62349486],"study_design_scores_gemma":[0.00020823408,0.00023158343,0.0026078667,0.00043859598,0.00017776823,0.0012010982,0.0030020585,0.30206168,0.03394187,0.3463539,0.30962026,0.00015523256],"about_ca_topic_score_codex":0.0027490635,"about_ca_topic_score_gemma":0.0037242237,"teacher_disagreement_score":0.008759454,"about_ca_system_score_codex":0.001169596,"about_ca_system_score_gemma":0.0019398745,"threshold_uncertainty_score":0.041430354},"labels":[],"label_agreement":null},{"id":"W2805005636","doi":"10.18653/v1/n18-2008","title":"Automatic Dialogue Generation with Expressed Emotions","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":145,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Computational linguistics; Natural language processing; Linguistics; Artificial intelligence; Volume (thermodynamics); Cognitive science; Psychology; Philosophy","score_opus":0.036668568292519306,"score_gpt":0.242900479940581,"score_spread":0.2062319116480617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805005636","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.103226736,0.0017415824,0.830764,0.001031847,0.0019289284,0.00073281414,0.0035029056,0.03353351,0.02353769],"genre_scores_gemma":[0.7163033,0.00033270408,0.26196703,0.00025166714,0.00040757348,0.0005972242,0.0058256285,0.0015425392,0.0127723105],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982791,0.0008883954,0.00006596811,0.0004393324,0.00020216721,0.0001250318],"domain_scores_gemma":[0.9981091,0.0010951805,0.000076651544,0.00016614418,0.00045845515,0.00009447148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012356379,0.0014467356,0.00071878795,0.0006815646,0.0005353385,0.0013976144,0.0009942761,0.0008481911,0.011368513],"category_scores_gemma":[0.005285103,0.00037220508,0.0007278406,0.0003070444,0.00030325662,0.0011202055,0.0020393527,0.00094795856,0.006614829],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002687716,0.00036505604,0.0027129375,0.001260112,0.00023965792,0.0010660167,0.00240485,0.014761279,0.10710963,0.008998776,0.073542565,0.78485143],"study_design_scores_gemma":[0.0003405038,0.0005657312,0.0036919273,0.0001936088,0.00024276435,0.0006936691,0.0016781974,0.8645619,0.06326611,0.02545731,0.03919683,0.00011145895],"about_ca_topic_score_codex":0.00041211492,"about_ca_topic_score_gemma":0.0005553169,"teacher_disagreement_score":0.011368513,"about_ca_system_score_codex":0.00034279146,"about_ca_system_score_gemma":0.00026732017,"threshold_uncertainty_score":0.03803146},"labels":[],"label_agreement":null},{"id":"W2805098445","doi":"","title":"TAC KBP 2015 : English Slot Filling Track Relational Learning with Expert Advice.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Advice (programming); Computer science; Artificial intelligence; Natural language processing; Programming language","score_opus":0.017555628212214344,"score_gpt":0.24894614480986596,"score_spread":0.2313905165976516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805098445","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046021555,0.0026703943,0.4019088,0.0033794534,0.0023696397,0.0013093002,0.24138065,0.23765375,0.06330642],"genre_scores_gemma":[0.16738752,0.0008315361,0.3299846,0.0010351487,0.0003710324,0.0012371251,0.44770384,0.011026392,0.040422834],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961216,0.0011602638,0.00040678764,0.00092092704,0.0011208793,0.00026953768],"domain_scores_gemma":[0.98697644,0.0051221796,0.00035551263,0.0034205066,0.0034979854,0.00062733004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004054413,0.001443813,0.0012380483,0.0027739399,0.0016247869,0.0038583535,0.005068196,0.0021876353,0.043043677],"category_scores_gemma":[0.03334397,0.0010704186,0.0007448547,0.002714881,0.0007932867,0.011063101,0.005094037,0.0028850287,0.034549735],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009162786,0.00035043978,0.002132191,0.0010810391,0.00010709057,0.00027807555,0.0009405496,0.003352538,0.0050600553,0.011568253,0.7380711,0.23614234],"study_design_scores_gemma":[0.0006958951,0.0003914276,0.0048531527,0.00063319836,0.00019016623,0.0005810136,0.0018797644,0.15684925,0.027322434,0.04954911,0.7567943,0.00026039744],"about_ca_topic_score_codex":0.02229614,"about_ca_topic_score_gemma":0.029731601,"teacher_disagreement_score":0.043043677,"about_ca_system_score_codex":0.0013504191,"about_ca_system_score_gemma":0.004321605,"threshold_uncertainty_score":0.1439954},"labels":[],"label_agreement":null},{"id":"W2805100503","doi":"","title":"Improve Neural Mention Detection and Classification via Enforced Training and Inference Consistency.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Consistency (knowledge bases); Computer science; Inference; Artificial intelligence; Machine learning; Training (meteorology); Training set; Pattern recognition (psychology)","score_opus":0.02995055309663209,"score_gpt":0.2787405406084506,"score_spread":0.24878998751181852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805100503","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04614866,0.0022800395,0.93596566,0.0017809751,0.0005002999,0.00013987561,0.00092176005,0.008933196,0.0033295914],"genre_scores_gemma":[0.6724212,0.00067836646,0.30937243,0.0014018958,0.0012347582,0.00024517873,0.0064036725,0.0010909451,0.007151506],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99436426,0.0021352577,0.00037501167,0.0015512125,0.0011584649,0.00041570442],"domain_scores_gemma":[0.9668299,0.021746727,0.0013274241,0.005949477,0.0034860084,0.00066043355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070880977,0.0013160105,0.0021557966,0.0026644673,0.0013008119,0.0025040943,0.005227287,0.0035981284,0.0041361516],"category_scores_gemma":[0.053984642,0.0008874874,0.0013779957,0.0026999125,0.0012414444,0.0064375536,0.00367275,0.0050735534,0.003063086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015447317,0.0007798393,0.012042848,0.0004774927,0.00061617256,0.00028636697,0.00044011665,0.11903685,0.009299121,0.02770668,0.045978956,0.78179085],"study_design_scores_gemma":[0.000054362186,0.00006837189,0.0011066548,0.000045346645,0.00008932807,0.00010760952,0.000050917115,0.95782965,0.00395311,0.034410533,0.0022635732,0.000020571688],"about_ca_topic_score_codex":0.006269614,"about_ca_topic_score_gemma":0.011670092,"teacher_disagreement_score":0.0070880977,"about_ca_system_score_codex":0.0012918704,"about_ca_system_score_gemma":0.00283182,"threshold_uncertainty_score":0.037485898},"labels":[],"label_agreement":null},{"id":"W2805220942","doi":"","title":"CMUML Micro-Reader System for KBP 2016 Cold Start Slot Filling, Event Nugget Detection, and Event Argument Linking.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Argument (complex analysis); Event (particle physics); Computer science; Physics; Astrophysics; Medicine","score_opus":0.010010295807845744,"score_gpt":0.23355224087906762,"score_spread":0.22354194507122188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805220942","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002799152,0.0005431931,0.26785737,0.00051624747,0.00046002524,0.00051307,0.048651654,0.65997124,0.018687977],"genre_scores_gemma":[0.0657671,0.00045810346,0.6592273,0.0011462158,0.0005469062,0.0026538863,0.13982937,0.0708878,0.059483375],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957801,0.0010128855,0.0005462367,0.0011650354,0.001242721,0.00025289436],"domain_scores_gemma":[0.98988557,0.004271255,0.0005603405,0.0024632607,0.0024008024,0.0004187425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00480649,0.0023301998,0.0016299085,0.0049750977,0.0011602648,0.004505567,0.0035612362,0.002356738,0.1459176],"category_scores_gemma":[0.024865974,0.0017584263,0.00090904563,0.0030510502,0.0005923519,0.006806406,0.004271916,0.0021701625,0.13578187],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016792632,0.00012886633,0.0018498135,0.0016969519,0.00012878832,0.00037821944,0.0010942958,0.0008848458,0.020154387,0.010674345,0.6761945,0.28513572],"study_design_scores_gemma":[0.00046491483,0.00021964626,0.002669081,0.00041267255,0.00009384183,0.000645679,0.00056110893,0.044492792,0.08339831,0.015161019,0.8515941,0.0002868754],"about_ca_topic_score_codex":0.003822357,"about_ca_topic_score_gemma":0.006401928,"teacher_disagreement_score":0.1459176,"about_ca_system_score_codex":0.0013364212,"about_ca_system_score_gemma":0.002456538,"threshold_uncertainty_score":0.48814297},"labels":[],"label_agreement":null},{"id":"W2805241394","doi":"","title":"Building Knowledge Bases with Universal Schema: Cold Start and Slot-Filling Approaches.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Schema (genetic algorithms); Computer science; Information retrieval","score_opus":0.03856681738170496,"score_gpt":0.24421428120514954,"score_spread":0.20564746382344457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805241394","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0072857337,0.00086742593,0.9853078,0.0003186751,0.00006744213,0.0002538096,0.0015685874,0.0025820981,0.0017484627],"genre_scores_gemma":[0.056313347,0.00059032935,0.93608075,0.00017688157,0.000046415553,0.00031840094,0.0049904445,0.000338453,0.0011450526],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99432325,0.0024115357,0.0006374032,0.0012487733,0.0010927598,0.00028638213],"domain_scores_gemma":[0.97728026,0.0130446255,0.00066737796,0.006136049,0.0024018264,0.0004699118],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008541439,0.0010089398,0.0017948938,0.0056602308,0.001840652,0.0054747006,0.004593896,0.001822442,0.0053987107],"category_scores_gemma":[0.035819393,0.0015614637,0.0033845408,0.0071234386,0.0019578235,0.014471432,0.0067287963,0.0028160883,0.0028173283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005139296,0.00041583323,0.0036756329,0.0022476967,0.00066227675,0.00058167777,0.004179651,0.026184166,0.0059154117,0.16497743,0.02979843,0.7608478],"study_design_scores_gemma":[0.00016060568,0.00022117743,0.0013915477,0.0011399236,0.0008575076,0.00082539424,0.0032312807,0.33405724,0.021507408,0.543681,0.09276173,0.00016515923],"about_ca_topic_score_codex":0.004838099,"about_ca_topic_score_gemma":0.00932504,"teacher_disagreement_score":0.008541439,"about_ca_system_score_codex":0.001108934,"about_ca_system_score_gemma":0.0037509724,"threshold_uncertainty_score":0.045171976},"labels":[],"label_agreement":null},{"id":"W2805275877","doi":"","title":"DoughnutPRIS at TAC KBP 2016.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.009630229271104514,"score_gpt":0.23452191201053615,"score_spread":0.22489168273943164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805275877","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00653349,0.010228225,0.13211718,0.10443477,0.046917543,0.000666184,0.09540281,0.05876566,0.5449342],"genre_scores_gemma":[0.025115423,0.0023187601,0.03735187,0.0035640139,0.004507925,0.00026894122,0.048668128,0.009548749,0.86865616],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963988,0.0009361169,0.00019698837,0.0008753782,0.0013134318,0.0002793158],"domain_scores_gemma":[0.9913488,0.0019395286,0.00026925298,0.0015281187,0.0035149949,0.0013994236],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0052322852,0.0014883867,0.0017617774,0.00376315,0.0035627766,0.0090893265,0.002784458,0.0022188365,0.43010634],"category_scores_gemma":[0.016928423,0.0007496204,0.0010286263,0.0035159194,0.0008857146,0.009632265,0.0043772655,0.0027975834,0.30597582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001432802,0.000051600302,0.00014873427,0.000109002554,0.000010684803,0.00011848025,0.00015045449,0.00021922465,0.0004725869,0.007994367,0.9434871,0.047094416],"study_design_scores_gemma":[0.00003642373,0.000013474598,0.0003178967,0.000094145085,0.000010994018,0.00005852423,0.00019825455,0.0018143323,0.00063116424,0.014568647,0.98222727,0.000028940143],"about_ca_topic_score_codex":0.019922335,"about_ca_topic_score_gemma":0.03745162,"teacher_disagreement_score":0.43010634,"about_ca_system_score_codex":0.003433547,"about_ca_system_score_gemma":0.0026702427,"threshold_uncertainty_score":0.81288415},"labels":[],"label_agreement":null},{"id":"W2805435982","doi":"","title":"The ZHI-EDL System for Entity Discovery and Linking at TAC KBP 2017.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.015520133304319911,"score_gpt":0.26132830922367434,"score_spread":0.24580817591935442,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805435982","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006298369,0.0008702967,0.3117216,0.0009140763,0.00038322975,0.00061464997,0.19245514,0.4687828,0.017959887],"genre_scores_gemma":[0.040168073,0.00065786275,0.35231033,0.0003849685,0.00015276489,0.0009808232,0.57166994,0.016614987,0.017060315],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978503,0.0006033304,0.0002952748,0.00048405598,0.0006218156,0.00014529521],"domain_scores_gemma":[0.995926,0.0014092913,0.00025111367,0.0014374695,0.00073635223,0.00023982284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045298357,0.0015089076,0.0012277225,0.0065926723,0.0011825778,0.003871584,0.0027674749,0.0013530196,0.04378251],"category_scores_gemma":[0.013034715,0.0011171029,0.0011289879,0.004194766,0.00051540177,0.007016392,0.005899436,0.0020560992,0.039128125],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001018543,0.00022409656,0.0036795565,0.0022237499,0.0003705416,0.00057019765,0.0012044274,0.0052809655,0.0076916446,0.02253974,0.6984357,0.25676078],"study_design_scores_gemma":[0.00051869417,0.00015084182,0.003919413,0.00038398072,0.00021891459,0.00040607303,0.00068245205,0.102608964,0.021684075,0.0416393,0.8275943,0.0001930964],"about_ca_topic_score_codex":0.009102065,"about_ca_topic_score_gemma":0.011158931,"teacher_disagreement_score":0.04378251,"about_ca_system_score_codex":0.0011316248,"about_ca_system_score_gemma":0.002660227,"threshold_uncertainty_score":0.14646709},"labels":[],"label_agreement":null},{"id":"W2805479698","doi":"","title":"HLTCOE Participation in TAC KBP 2017: Cold Start and EDL Pilot.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Aeronautics; Engineering","score_opus":0.030261234122667523,"score_gpt":0.30120395317847226,"score_spread":0.27094271905580475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805479698","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60120434,0.0011375379,0.11197787,0.013228368,0.0044911443,0.007346865,0.06185928,0.07854726,0.12020737],"genre_scores_gemma":[0.7219363,0.00035142055,0.08431195,0.0022353877,0.0006583018,0.006855862,0.10535894,0.00888017,0.06941171],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98750985,0.0068086786,0.0006353938,0.0011381435,0.0029387234,0.00096915645],"domain_scores_gemma":[0.9182078,0.04570004,0.0009107219,0.009615849,0.019271474,0.006294199],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01659854,0.00071625045,0.0007454118,0.0013995824,0.002319814,0.003069542,0.0020186256,0.0016906479,0.023201296],"category_scores_gemma":[0.074021555,0.00050873397,0.00046831756,0.0013343527,0.00094573305,0.0056168344,0.005229705,0.003239181,0.010532202],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0058631934,0.004401803,0.011793832,0.0021754126,0.00013223385,0.0012447357,0.038587105,0.0041931495,0.015007943,0.0047655627,0.61886996,0.29296508],"study_design_scores_gemma":[0.0022740345,0.002591193,0.021513617,0.000894646,0.00020234383,0.0004390068,0.032874625,0.04504442,0.026047159,0.009806321,0.85787344,0.0004392958],"about_ca_topic_score_codex":0.020700887,"about_ca_topic_score_gemma":0.03355832,"teacher_disagreement_score":0.023201296,"about_ca_system_score_codex":0.0015488148,"about_ca_system_score_gemma":0.005149947,"threshold_uncertainty_score":0.0877825},"labels":[],"label_agreement":null},{"id":"W2805547232","doi":"","title":"WIP Event Detection System at TAC KBP 2016 Event Nugget Track.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Event (particle physics); Computer science; Real-time computing; Operating system; Physics","score_opus":0.007172369900705897,"score_gpt":0.2281875531371292,"score_spread":0.22101518323642333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805547232","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07158066,0.0020604015,0.26334992,0.001768051,0.0012907974,0.0014869339,0.25700614,0.35070336,0.050753802],"genre_scores_gemma":[0.27183643,0.00075281,0.1687491,0.00061571255,0.0003948844,0.0012779232,0.513911,0.0065204636,0.035941783],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987338,0.00014858582,0.00010850635,0.0003435271,0.0005582585,0.000107349355],"domain_scores_gemma":[0.99730957,0.0005105834,0.00021471428,0.0004497436,0.0012752758,0.00024006303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018051389,0.0009736175,0.0010272514,0.0036865128,0.00087498024,0.0017125192,0.0016611804,0.00096378464,0.016557036],"category_scores_gemma":[0.0071042953,0.00046187433,0.00033180235,0.0021336176,0.00018759645,0.0030149121,0.0014148604,0.0012455727,0.017184623],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014632933,0.00022766428,0.01195475,0.0009527626,0.00017847227,0.00044369523,0.0005381944,0.004760416,0.013831176,0.0022424124,0.7187768,0.24463032],"study_design_scores_gemma":[0.000554903,0.00055072387,0.03041633,0.00024232255,0.00027799435,0.0005292912,0.0010816256,0.2797695,0.053142767,0.012433554,0.6207959,0.0002050431],"about_ca_topic_score_codex":0.012173267,"about_ca_topic_score_gemma":0.012884902,"teacher_disagreement_score":0.016557036,"about_ca_system_score_codex":0.00066476205,"about_ca_system_score_gemma":0.0014563958,"threshold_uncertainty_score":0.05538881},"labels":[],"label_agreement":null},{"id":"W2805550868","doi":"","title":"New York University 2016 System for KBP Event Nugget: A Deep Learning Approach.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Computer science; Artificial intelligence; History; Physics","score_opus":0.011517172097703287,"score_gpt":0.20961021764074636,"score_spread":0.19809304554304308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805550868","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025639208,0.0015736365,0.17969595,0.0019886692,0.0012672495,0.0010144161,0.23859344,0.5115215,0.038706],"genre_scores_gemma":[0.15542711,0.0010406073,0.25775772,0.00076426606,0.00038414044,0.0012429314,0.5027919,0.013895547,0.06669586],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993531,0.000104594736,0.00006438667,0.0002341434,0.0001756691,0.000068092],"domain_scores_gemma":[0.99839634,0.00041398435,0.000086368724,0.00048617378,0.00044365737,0.00017343462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013536086,0.0011802511,0.0010769899,0.002675449,0.00075908436,0.0017729711,0.002011681,0.0012192909,0.04720232],"category_scores_gemma":[0.007068744,0.0006810449,0.0006208201,0.001847621,0.00028691726,0.0045431433,0.0028891838,0.0019544628,0.032268915],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011153896,0.00022724466,0.0036414803,0.00062377454,0.00015985426,0.0002513139,0.00035568018,0.005111031,0.004567222,0.0061995527,0.759908,0.21783943],"study_design_scores_gemma":[0.0005481301,0.00024850894,0.006497366,0.00024703625,0.00017302191,0.0003923448,0.00030737356,0.3394852,0.023430426,0.027059698,0.60141474,0.00019618557],"about_ca_topic_score_codex":0.019251307,"about_ca_topic_score_gemma":0.024890602,"teacher_disagreement_score":0.04720232,"about_ca_system_score_codex":0.0010955756,"about_ca_system_score_gemma":0.0018895,"threshold_uncertainty_score":0.15790749},"labels":[],"label_agreement":null},{"id":"W2805588680","doi":"","title":"Multi-lingual Extraction and Integration of Entities, Relations, Events and Sentiments into ColdStart++ KBs with the SAFT System.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Extraction (chemistry); Natural language processing; Information retrieval; Artificial intelligence; Chromatography; Chemistry","score_opus":0.012792609573497168,"score_gpt":0.2729504194354924,"score_spread":0.26015780986199527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805588680","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026749369,0.0015078317,0.53067863,0.0013352846,0.00082373305,0.0012407781,0.27292594,0.14280997,0.021928504],"genre_scores_gemma":[0.060841694,0.0006358948,0.48957944,0.0003292843,0.00019905949,0.0007464706,0.43154562,0.006563396,0.00955914],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981207,0.00042295962,0.000330078,0.0006162301,0.00040707103,0.00010296275],"domain_scores_gemma":[0.9952592,0.0017284588,0.00030803488,0.0008152733,0.0016965218,0.00019261408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017847857,0.0022372697,0.0014562731,0.0056041447,0.0015548088,0.0033685106,0.0016634626,0.0012679375,0.019202096],"category_scores_gemma":[0.008293246,0.0013845921,0.0018579633,0.0043437737,0.00058839424,0.005806079,0.0040294416,0.002227919,0.026027823],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008452745,0.00033115508,0.0078012426,0.0048163617,0.000603839,0.0021267007,0.004025173,0.004142345,0.04326824,0.014652704,0.46545094,0.4519361],"study_design_scores_gemma":[0.00026133202,0.00028535223,0.018120237,0.0011950934,0.0007651709,0.0022223233,0.0041492837,0.11689741,0.0454298,0.040078748,0.77018636,0.0004089449],"about_ca_topic_score_codex":0.010154069,"about_ca_topic_score_gemma":0.021267382,"teacher_disagreement_score":0.019202096,"about_ca_system_score_codex":0.00094621617,"about_ca_system_score_gemma":0.0022952098,"threshold_uncertainty_score":0.064237356},"labels":[],"label_agreement":null},{"id":"W2805607905","doi":"","title":"UZH at TAC KBP 2017: Event Nugget Detection via Joint Learning with Softmax-Margin Objective.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Softmax function; Margin (machine learning); Joint (building); Computer science; Event (particle physics); Artificial intelligence; Machine learning; Deep learning; Engineering; Physics; Structural engineering","score_opus":0.010434209432448077,"score_gpt":0.23880144636835732,"score_spread":0.22836723693590924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805607905","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037653685,0.0077164946,0.85071826,0.013588669,0.012010786,0.0005418775,0.02014104,0.04246541,0.0151638165],"genre_scores_gemma":[0.31251228,0.0017842754,0.54733706,0.0033116557,0.0049055717,0.00080397405,0.06484639,0.007081206,0.05741751],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99762565,0.0009807833,0.00009193993,0.0004857999,0.0005908185,0.00022497836],"domain_scores_gemma":[0.99579865,0.0023408353,0.00008168389,0.00062087894,0.00073186203,0.00042614312],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051452317,0.002965295,0.0039591473,0.0018454415,0.0014771158,0.0034032625,0.0029145386,0.0038583402,0.0155191785],"category_scores_gemma":[0.016437665,0.0011963431,0.0015563249,0.0023030206,0.00114594,0.0046687084,0.0042428076,0.006792116,0.0099875685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024451741,0.0004340973,0.0010251744,0.00044012166,0.00041623728,0.0004288646,0.00027231517,0.055541094,0.0042941486,0.013119589,0.6643685,0.2572147],"study_design_scores_gemma":[0.00044208803,0.00015922212,0.0007825227,0.000065181935,0.00007901484,0.00012400563,0.00013882053,0.9054771,0.00482179,0.053656217,0.03418095,0.000073100135],"about_ca_topic_score_codex":0.013725979,"about_ca_topic_score_gemma":0.018036272,"teacher_disagreement_score":0.0155191785,"about_ca_system_score_codex":0.0013617073,"about_ca_system_score_gemma":0.0021525514,"threshold_uncertainty_score":0.051916778},"labels":[],"label_agreement":null},{"id":"W2805676314","doi":"10.18653/v1/s18-1168","title":"UNBNLP at SemEval-2018 Task 10: Evaluating unsupervised approaches to capturing discriminative attributes","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Discriminative model; Computer science; SemEval; WordNet; Artificial intelligence; Word (group theory); Sentence; Natural language processing; Task (project management); Simple (philosophy); Machine learning; Mathematics","score_opus":0.31558798624108264,"score_gpt":0.3145837265306022,"score_spread":0.0010042597104804596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805676314","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.674235,0.015930826,0.124545425,0.0025295748,0.0043366957,0.005286382,0.08482553,0.05070053,0.037609987],"genre_scores_gemma":[0.51827496,0.0018608157,0.12164122,0.0020092565,0.0011007424,0.003959207,0.32922837,0.0032490285,0.0186765],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99023813,0.0046767746,0.00069018424,0.0024084616,0.0014501885,0.000536249],"domain_scores_gemma":[0.9789979,0.011061847,0.0010930027,0.0048177615,0.002457322,0.0015721672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012793632,0.0064528487,0.0034129636,0.0037233601,0.0021180091,0.0035034684,0.005760831,0.0061984276,0.0075803036],"category_scores_gemma":[0.021482175,0.0008733614,0.0025874968,0.0028628553,0.001711515,0.007983363,0.005320428,0.005372645,0.009631955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0091348775,0.012827731,0.0403855,0.0041283527,0.0031023177,0.000976304,0.00071135245,0.062050667,0.01582259,0.0031601388,0.2663923,0.5813079],"study_design_scores_gemma":[0.002827946,0.008383955,0.0555551,0.00079231156,0.0015553384,0.0032077192,0.0020153737,0.79329497,0.044260677,0.014245216,0.073187366,0.0006740749],"about_ca_topic_score_codex":0.010452511,"about_ca_topic_score_gemma":0.018254632,"teacher_disagreement_score":0.012793632,"about_ca_system_score_codex":0.0018719424,"about_ca_system_score_gemma":0.0022044878,"threshold_uncertainty_score":0.067659974},"labels":[],"label_agreement":null},{"id":"W2805707541","doi":"","title":"Overview of SYDNEY System for TAC KBP 2015 Event Nugget Detection.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Computer science; Real-time computing; Physics","score_opus":0.031325365726029175,"score_gpt":0.2926198252896208,"score_spread":0.26129445956359165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805707541","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019600017,0.0032308553,0.40491837,0.0012165633,0.00073208,0.002701121,0.052925196,0.47011322,0.04456251],"genre_scores_gemma":[0.12716082,0.0019061167,0.60011107,0.0011165752,0.0005204212,0.0028336209,0.1879775,0.014194985,0.06417888],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976562,0.0003830345,0.000248446,0.00063065527,0.0008987205,0.00018286149],"domain_scores_gemma":[0.99691415,0.0005255474,0.00015202862,0.0006592247,0.0014825634,0.00026652456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024137017,0.001342619,0.0015390532,0.007785419,0.0012049507,0.002982629,0.002451079,0.0011513409,0.020248579],"category_scores_gemma":[0.0077412394,0.0011014574,0.0010663398,0.003157192,0.000464227,0.003939849,0.0026129757,0.0012818138,0.03394531],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007282428,0.00027390628,0.0074748993,0.0013737363,0.00039736863,0.00041983364,0.0013462378,0.0041890657,0.019968245,0.005060781,0.5571268,0.40164095],"study_design_scores_gemma":[0.00021156685,0.00031484175,0.0141046075,0.0003798901,0.00037935283,0.00088830595,0.0010055458,0.17349312,0.03674164,0.014892427,0.7572351,0.00035365284],"about_ca_topic_score_codex":0.024625344,"about_ca_topic_score_gemma":0.030275375,"teacher_disagreement_score":0.024625344,"about_ca_system_score_codex":0.001132508,"about_ca_system_score_gemma":0.0033177526,"threshold_uncertainty_score":0.067738295},"labels":[],"label_agreement":null},{"id":"W2805727849","doi":"","title":"UI CCG TAC-KBP2017 Submissions: Entity Discovery and Linking, and Event Nugget Detection and Co-reference.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Computer science; Physics","score_opus":0.016664417114219225,"score_gpt":0.2829101230503985,"score_spread":0.26624570593617924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805727849","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010929875,0.0025641103,0.07054429,0.013299456,0.03963488,0.002963577,0.671537,0.10350188,0.08502488],"genre_scores_gemma":[0.013186518,0.0005027541,0.05054872,0.001173062,0.0040241764,0.0013807622,0.83469677,0.01834765,0.07613961],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9876923,0.0037627253,0.0009102752,0.0014236975,0.0049062837,0.0013046548],"domain_scores_gemma":[0.9409206,0.01348983,0.0014270599,0.012141072,0.022431947,0.009589451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013128591,0.0035657552,0.00303747,0.006278628,0.003922259,0.008118222,0.005334401,0.0051360065,0.21414426],"category_scores_gemma":[0.06114899,0.001330785,0.0021682798,0.00544526,0.000992148,0.006000506,0.008833889,0.0034207106,0.17188972],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019856587,0.000051205534,0.00016140072,0.00018155428,0.000016927328,0.000078742174,0.00003687458,0.00023191114,0.00041016418,0.00052048033,0.9902174,0.0078948205],"study_design_scores_gemma":[0.00039390483,0.0001332711,0.0025319888,0.0002188283,0.000043920485,0.00036933774,0.00029433594,0.009050394,0.0036306991,0.006655429,0.97657245,0.00010553847],"about_ca_topic_score_codex":0.020978216,"about_ca_topic_score_gemma":0.048482474,"teacher_disagreement_score":0.21414426,"about_ca_system_score_codex":0.002701175,"about_ca_system_score_gemma":0.0072922874,"threshold_uncertainty_score":0.71638393},"labels":[],"label_agreement":null},{"id":"W2805736953","doi":"","title":"REDES at TAC Knowledge Base Population 2016 : EDL and BeSt tracks.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Base (topology); Computer science; Knowledge base; Population; Artificial intelligence; Mathematics; Demography; Mathematical analysis","score_opus":0.013606234031065272,"score_gpt":0.2543820827541767,"score_spread":0.2407758487231114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805736953","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.075915955,0.005932938,0.44412547,0.011264145,0.003064578,0.0011545478,0.22080094,0.13353035,0.10421109],"genre_scores_gemma":[0.19809562,0.0017312437,0.38152638,0.0012357144,0.00050157023,0.0007498662,0.33837467,0.011260793,0.06652418],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9955552,0.0012349793,0.00038019847,0.0009823546,0.0015175039,0.00032962736],"domain_scores_gemma":[0.97881615,0.0057956576,0.00059637136,0.007992377,0.0052869925,0.001512525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008540756,0.000817292,0.0009215069,0.006207281,0.0017031571,0.0061137835,0.0029344787,0.0016733245,0.025693616],"category_scores_gemma":[0.04549226,0.0009343388,0.0010339121,0.0061201523,0.0006008009,0.01326908,0.0042639947,0.0026010764,0.01619013],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005613433,0.00038918335,0.008560446,0.00047302063,0.00012131473,0.000104655264,0.0007163316,0.008541968,0.0014773939,0.018201642,0.4376863,0.5231664],"study_design_scores_gemma":[0.00028807123,0.000287207,0.0067692967,0.00047976716,0.00018590117,0.00037768474,0.0011105251,0.14897105,0.009354836,0.065042295,0.7669983,0.00013511861],"about_ca_topic_score_codex":0.01885472,"about_ca_topic_score_gemma":0.03632708,"teacher_disagreement_score":0.025693616,"about_ca_system_score_codex":0.0023059403,"about_ca_system_score_gemma":0.0054647634,"threshold_uncertainty_score":0.08595371},"labels":[],"label_agreement":null},{"id":"W2805757823","doi":"","title":"JHU/APL CONSENSE: The Confidence-Based System Ensembler.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.025571649878529317,"score_gpt":0.2540874549794681,"score_spread":0.2285158051009388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805757823","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0075419205,0.00093647017,0.9266001,0.0004817107,0.00045198965,0.00021318698,0.003568721,0.053477343,0.0067285304],"genre_scores_gemma":[0.21163727,0.00066108536,0.7501421,0.00051477674,0.0005283901,0.0005425663,0.016018156,0.0044782492,0.0154773025],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976019,0.00067389675,0.00010312104,0.00050507614,0.0009926303,0.00012330116],"domain_scores_gemma":[0.9961132,0.0010561,0.0001510216,0.0013370697,0.001163905,0.00017867525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030857038,0.0012511547,0.0014631002,0.0019019716,0.00081896264,0.0017588516,0.002393552,0.0014201386,0.011796361],"category_scores_gemma":[0.015515611,0.00057951256,0.0007256859,0.0016346906,0.0004174845,0.003290326,0.003358827,0.0023068169,0.008240406],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055710215,0.00021795306,0.002261792,0.00044222543,0.00038999153,0.00018392982,0.00019466603,0.04506762,0.010452028,0.01458536,0.1412608,0.7843865],"study_design_scores_gemma":[0.00007230025,0.00013730269,0.0010550736,0.000043763914,0.0001006199,0.00024270183,0.000045401368,0.9222878,0.016023526,0.020797627,0.03912968,0.00006424694],"about_ca_topic_score_codex":0.0033931208,"about_ca_topic_score_gemma":0.0046870424,"teacher_disagreement_score":0.011796361,"about_ca_system_score_codex":0.0004991507,"about_ca_system_score_gemma":0.0012690139,"threshold_uncertainty_score":0.039462745},"labels":[],"label_agreement":null},{"id":"W2805760825","doi":"","title":"CRIM’s Systems for the Tri-lingual Entity Detection and Linking Task.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Task (project management); Natural language processing; Engineering; Systems engineering","score_opus":0.021790547674589815,"score_gpt":0.27742051027554476,"score_spread":0.25562996260095494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805760825","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031449128,0.0047576292,0.58144623,0.002577594,0.0019376022,0.0014397965,0.038007576,0.27104473,0.067339726],"genre_scores_gemma":[0.16460961,0.0017912003,0.7117902,0.0014658602,0.0006057842,0.0015947812,0.06351745,0.0100598,0.04456526],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971998,0.00068586116,0.00020860312,0.0009365264,0.0006973087,0.00027192655],"domain_scores_gemma":[0.99333173,0.0019587541,0.0002917109,0.0028172145,0.0013437525,0.00025683266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035346297,0.0017309082,0.0009644452,0.005832145,0.0018115607,0.0028971047,0.0032662086,0.002511596,0.027553402],"category_scores_gemma":[0.01702822,0.0010991524,0.0013098377,0.005259942,0.00057070475,0.006991453,0.0051944735,0.002248107,0.026297837],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038786748,0.00029732788,0.0041591376,0.0006258249,0.00032540417,0.00026744648,0.0005871382,0.0028875011,0.0056563728,0.014492458,0.4236624,0.54665107],"study_design_scores_gemma":[0.00019792002,0.00025782466,0.007895945,0.0002259143,0.00042448103,0.0016435481,0.0006911428,0.21534303,0.046006583,0.04635772,0.6807221,0.00023373737],"about_ca_topic_score_codex":0.010385001,"about_ca_topic_score_gemma":0.019858427,"teacher_disagreement_score":0.027553402,"about_ca_system_score_codex":0.0011576683,"about_ca_system_score_gemma":0.0029216476,"threshold_uncertainty_score":0.092175364},"labels":[],"label_agreement":null},{"id":"W2805764238","doi":"10.63317/2qe447uqn7co","title":"Revisiting the Task of Scoring Open IE Relations","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Task (project management); Engineering; Systems engineering","score_opus":0.05977835686829329,"score_gpt":0.30512455183047277,"score_spread":0.24534619496217946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805764238","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07369484,0.0051007275,0.8737772,0.010496731,0.0018322743,0.00040479598,0.0045101424,0.00627768,0.023905672],"genre_scores_gemma":[0.5453355,0.003753714,0.40803248,0.0019379386,0.003862507,0.00036952255,0.013908851,0.0026756793,0.020123824],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9866431,0.0062686945,0.001204458,0.0023518244,0.002799076,0.0007328846],"domain_scores_gemma":[0.93818796,0.04254129,0.002073143,0.008229719,0.00738311,0.0015846111],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01066073,0.0017301579,0.0020431934,0.005957642,0.0028091818,0.007994417,0.0027537383,0.003801371,0.008648866],"category_scores_gemma":[0.08070514,0.0008941244,0.001796652,0.0073973113,0.0017505885,0.012311235,0.006973114,0.0073319874,0.007214076],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055984216,0.0004919214,0.01748179,0.0014322707,0.00029747485,0.00063955673,0.0026748595,0.010950911,0.01436169,0.066263825,0.08920532,0.79564047],"study_design_scores_gemma":[0.00016564543,0.00036589877,0.022769658,0.000974878,0.00053070375,0.0026995256,0.0053398074,0.4238771,0.026541471,0.37973818,0.13673536,0.00026180348],"about_ca_topic_score_codex":0.005975315,"about_ca_topic_score_gemma":0.007862786,"teacher_disagreement_score":0.01066073,"about_ca_system_score_codex":0.001034386,"about_ca_system_score_gemma":0.0027000075,"threshold_uncertainty_score":0.056379974},"labels":[],"label_agreement":null},{"id":"W2805780438","doi":"","title":"GWU English TAC-KBP EL Diagnostic Task with Name Mention.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Task (project management); Computer science; Natural language processing; Linguistics; Philosophy; Economics; Management","score_opus":0.009584582317296956,"score_gpt":0.23523312238330743,"score_spread":0.22564854006601048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805780438","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15353984,0.0038122074,0.1711615,0.008369769,0.003690178,0.0016765532,0.38003054,0.105938226,0.17178126],"genre_scores_gemma":[0.43260193,0.0006336249,0.09912124,0.002260528,0.00050910015,0.00082623924,0.42106032,0.0041790563,0.03880797],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982401,0.000433644,0.00019552269,0.00057515677,0.00035100285,0.00020453078],"domain_scores_gemma":[0.99332327,0.0040048044,0.0002192934,0.00076369935,0.0014212402,0.00026774567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013315165,0.0018405626,0.0010825719,0.0025523833,0.0015147403,0.0020820838,0.0018196321,0.0028325324,0.049141567],"category_scores_gemma":[0.013289085,0.0005071599,0.0007839088,0.0019260993,0.0005201257,0.0059950138,0.0032233605,0.0018881939,0.031450707],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011752002,0.00028726962,0.0038406758,0.0027140998,0.00012640063,0.0035013182,0.00082955067,0.003618658,0.011324469,0.0060963845,0.78129977,0.18518618],"study_design_scores_gemma":[0.00082156673,0.0003597943,0.017359935,0.0012689832,0.00037398952,0.008742222,0.0067459694,0.14808486,0.06966134,0.0506416,0.69561523,0.0003245118],"about_ca_topic_score_codex":0.010144451,"about_ca_topic_score_gemma":0.011039435,"teacher_disagreement_score":0.049141567,"about_ca_system_score_codex":0.0009970141,"about_ca_system_score_gemma":0.0016244225,"threshold_uncertainty_score":0.16439492},"labels":[],"label_agreement":null},{"id":"W2805879805","doi":"","title":"The 2016 TAC KBP BeSt Evaluation.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.01721926072958735,"score_gpt":0.2739869869615991,"score_spread":0.2567677262320118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805879805","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04685632,0.020823808,0.049381923,0.04104441,0.023383064,0.0028856164,0.4751851,0.050841976,0.28959778],"genre_scores_gemma":[0.07981133,0.0030184302,0.063131906,0.0041522575,0.001954533,0.0029895557,0.7390485,0.010017769,0.095875755],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97796625,0.008984576,0.0015191961,0.002154963,0.008255203,0.0011198063],"domain_scores_gemma":[0.93798727,0.0149320215,0.001420028,0.00858861,0.031619042,0.005452952],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020120127,0.0029618389,0.002306683,0.008467178,0.0043819747,0.010756076,0.00528188,0.0036486431,0.048070524],"category_scores_gemma":[0.09291692,0.0012569083,0.0015034927,0.0059200223,0.0018720913,0.012980598,0.0071355426,0.0042759567,0.053905383],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006288141,0.00019855313,0.00069883297,0.00066146546,0.000066979315,0.00008192293,0.00018016643,0.0010874259,0.0006935309,0.0021668987,0.9380705,0.055465],"study_design_scores_gemma":[0.0016513118,0.0004469059,0.0077733076,0.0016965992,0.00033737827,0.00047842044,0.0012582247,0.020902203,0.005808852,0.014722691,0.94465786,0.0002662562],"about_ca_topic_score_codex":0.06860908,"about_ca_topic_score_gemma":0.07422434,"teacher_disagreement_score":0.06860908,"about_ca_system_score_codex":0.0049338792,"about_ca_system_score_gemma":0.009554579,"threshold_uncertainty_score":0.1608119},"labels":[],"label_agreement":null},{"id":"W2805986796","doi":"","title":"BUPT-PRIS System for TAC 2017 Event Nugget Detection, Event Argument Linking and ADR Tracks.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Argument (complex analysis); Event (particle physics); Computer science; Event data; Real-time computing; Medicine; Physics; Internal medicine","score_opus":0.013931735514242951,"score_gpt":0.26943956130827224,"score_spread":0.2555078257940293,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805986796","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009389588,0.0006858165,0.13296594,0.00040917975,0.0005814796,0.00064088235,0.134122,0.6946368,0.026568271],"genre_scores_gemma":[0.09978989,0.00056868093,0.29474795,0.00075980375,0.0006446839,0.0020270564,0.51039284,0.037741352,0.053327765],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988933,0.0001603529,0.00012508678,0.00037394976,0.0003499354,0.000097347096],"domain_scores_gemma":[0.9976803,0.0007016708,0.00021511369,0.00064116623,0.00059114734,0.0001705444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015630651,0.0016283466,0.0012558211,0.0044030403,0.0007239008,0.0027849136,0.0019090672,0.001527653,0.08426894],"category_scores_gemma":[0.0076845055,0.00071134797,0.001069022,0.0019916839,0.00034099372,0.004048262,0.0030326925,0.0013748967,0.07377339],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011455063,0.00016539088,0.005222846,0.0010837651,0.00025147578,0.000471533,0.0006587007,0.0017000706,0.008071928,0.0048423437,0.7763961,0.1999904],"study_design_scores_gemma":[0.0004948622,0.0002639587,0.01257696,0.00034857407,0.00028508613,0.0010086902,0.0007565577,0.12154709,0.036427267,0.022030098,0.8039922,0.0002685876],"about_ca_topic_score_codex":0.0036126834,"about_ca_topic_score_gemma":0.006198328,"teacher_disagreement_score":0.08426894,"about_ca_system_score_codex":0.00063039816,"about_ca_system_score_gemma":0.001237526,"threshold_uncertainty_score":0.28190768},"labels":[],"label_agreement":null},{"id":"W2806045113","doi":"","title":"University of Washington TAC-KBP 2016 System Description.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.010042888629884296,"score_gpt":0.19710016967162114,"score_spread":0.18705728104173686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806045113","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007198584,0.0019821166,0.20929621,0.003131645,0.000871933,0.0014536957,0.43069655,0.24346463,0.10190456],"genre_scores_gemma":[0.04586046,0.0015927017,0.10657157,0.0010945915,0.00016915063,0.0020572403,0.7767974,0.014092876,0.05176392],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979621,0.00042199038,0.00033989144,0.00038807234,0.0007317916,0.00015614227],"domain_scores_gemma":[0.9955663,0.00075336057,0.00022565054,0.0013245314,0.0018711523,0.00025909653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002460261,0.0014350292,0.0010921677,0.0034960834,0.0011511254,0.004866313,0.0032397457,0.0015418915,0.044843793],"category_scores_gemma":[0.009991113,0.0009764257,0.0008932447,0.0035039044,0.00040177585,0.0053110365,0.0018675632,0.002358756,0.054057527],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043406626,0.00010980463,0.001551486,0.0013735984,0.000076660814,0.00023983672,0.00034866415,0.004948527,0.0034969647,0.015963268,0.9033499,0.06810728],"study_design_scores_gemma":[0.00014729927,0.000065019536,0.0012832971,0.00033605393,0.00008091225,0.00037117797,0.00013872831,0.03455438,0.0077711293,0.012192739,0.94296634,0.00009293159],"about_ca_topic_score_codex":0.03597101,"about_ca_topic_score_gemma":0.02914614,"teacher_disagreement_score":0.044843793,"about_ca_system_score_codex":0.001828165,"about_ca_system_score_gemma":0.0045534717,"threshold_uncertainty_score":0.15001744},"labels":[],"label_agreement":null},{"id":"W2806067301","doi":"","title":"Target Focused Sentiment Extraction Framework.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Extraction (chemistry); Computer science; Sentiment analysis; Data science; Natural language processing; Chemistry; Chromatography","score_opus":0.015032931882100175,"score_gpt":0.2846347443300141,"score_spread":0.2696018124479139,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806067301","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015490353,0.0029750469,0.91078216,0.0017868547,0.0011373576,0.0011010561,0.015313909,0.010535253,0.040878],"genre_scores_gemma":[0.2631286,0.0025583094,0.62961215,0.0011994728,0.0013650446,0.0018322301,0.050297838,0.0013589502,0.04864741],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988819,0.00029145577,0.00009376325,0.00028073168,0.0003453052,0.00010695128],"domain_scores_gemma":[0.99880445,0.0003271675,0.000097593125,0.00014464887,0.00056068116,0.00006542489],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015406427,0.0012031919,0.00086681027,0.0036092761,0.0008205633,0.0022571168,0.0009927659,0.0011552254,0.009660386],"category_scores_gemma":[0.0041263327,0.0003811616,0.0010698253,0.0027551227,0.0003071943,0.0028473227,0.0017164022,0.0014234532,0.012847536],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036160104,0.000247737,0.0034003295,0.0012334206,0.00025874472,0.00032011067,0.0009907185,0.0024912788,0.03855663,0.044304658,0.21726316,0.6905716],"study_design_scores_gemma":[0.0001243542,0.0003118335,0.012037414,0.00049474667,0.0004954662,0.0013299999,0.0014500264,0.29124883,0.048995763,0.11341441,0.5299545,0.00014275864],"about_ca_topic_score_codex":0.0019522188,"about_ca_topic_score_gemma":0.0037515354,"teacher_disagreement_score":0.009660386,"about_ca_system_score_codex":0.0007002471,"about_ca_system_score_gemma":0.0011839748,"threshold_uncertainty_score":0.03231722},"labels":[],"label_agreement":null},{"id":"W2806148218","doi":"","title":"CMU-LTI at KBP 2015 Event Track.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Computer science; Event (particle physics); Physics; Operating system","score_opus":0.020276339394370485,"score_gpt":0.27516280170773244,"score_spread":0.254886462313362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806148218","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0039309864,0.0025136294,0.06665031,0.01817061,0.011750513,0.00093541836,0.589245,0.15597357,0.15082988],"genre_scores_gemma":[0.021095166,0.0011053303,0.03105728,0.0021236097,0.0029269112,0.0007763875,0.7745776,0.018947313,0.14739043],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957753,0.0011295824,0.00021532754,0.0008435758,0.0015194397,0.00051680306],"domain_scores_gemma":[0.98915625,0.0025981443,0.00036670276,0.002176926,0.0032067322,0.002495201],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.007046612,0.0030586438,0.0028715176,0.00330327,0.0025824818,0.009046711,0.0043204664,0.003498887,0.39681014],"category_scores_gemma":[0.026896026,0.0010872751,0.0011767244,0.005089399,0.00057084544,0.010084026,0.0052136225,0.004467943,0.38548568],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015032085,0.000029444462,0.0001145386,0.000101620455,0.0000069343982,0.000024093231,0.000025961139,0.00018418959,0.0002608121,0.0020352174,0.9883569,0.008710041],"study_design_scores_gemma":[0.0002294836,0.000122056605,0.001158802,0.00016044664,0.000019901016,0.000082673054,0.0001199813,0.008975287,0.002422306,0.010415699,0.9762314,0.000062042265],"about_ca_topic_score_codex":0.017620213,"about_ca_topic_score_gemma":0.01873323,"teacher_disagreement_score":0.39681014,"about_ca_system_score_codex":0.0030947602,"about_ca_system_score_gemma":0.003054455,"threshold_uncertainty_score":0.86037713},"labels":[],"label_agreement":null},{"id":"W2806185932","doi":"","title":"Cross Lingual Mention and Entity Embeddings for Cross-Lingual Entity Disambiguation.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Entity linking; Natural language processing; Artificial intelligence; Information retrieval; Knowledge base","score_opus":0.01332485124123976,"score_gpt":0.31413570805798635,"score_spread":0.3008108568167466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806185932","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015399064,0.004930483,0.961642,0.0010221372,0.0009643451,0.00015163964,0.004774647,0.0058578104,0.0052579176],"genre_scores_gemma":[0.27390122,0.0025951893,0.6831229,0.0005980425,0.0007083961,0.00043765697,0.029725319,0.0011010006,0.0078102835],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937907,0.0027621791,0.0007437988,0.0015443274,0.0007459118,0.00041319733],"domain_scores_gemma":[0.98822045,0.004845072,0.00060074014,0.0038709228,0.0021041937,0.00035866155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061136247,0.0011980799,0.001470324,0.008795739,0.0019459521,0.0037547767,0.002887046,0.00248376,0.006861695],"category_scores_gemma":[0.022949426,0.0010148571,0.0016845361,0.010838046,0.0010245123,0.012660031,0.0067917104,0.0030727352,0.007163097],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042080495,0.00033734887,0.0061249156,0.00094115845,0.00057528296,0.00044283856,0.0011806287,0.014825239,0.0050634537,0.09429086,0.08740999,0.7883874],"study_design_scores_gemma":[0.000055072218,0.000119580436,0.004768186,0.00036295154,0.00035269707,0.0008143669,0.0018835937,0.52377874,0.010443377,0.3665964,0.09067937,0.00014561649],"about_ca_topic_score_codex":0.0052643223,"about_ca_topic_score_gemma":0.01100486,"teacher_disagreement_score":0.008795739,"about_ca_system_score_codex":0.00088383676,"about_ca_system_score_gemma":0.0023685305,"threshold_uncertainty_score":0.0323323},"labels":[],"label_agreement":null},{"id":"W2806206036","doi":"","title":"UTD’s Event Nugget Detection and Coreference System at KBP 2015.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Coreference; Computer science; Event (particle physics); Natural language processing; Artificial intelligence; Resolution (logic); Physics; Astrophysics","score_opus":0.018472210104049498,"score_gpt":0.2522096874338165,"score_spread":0.233737477329767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806206036","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10996925,0.0044508236,0.31231755,0.0035816107,0.003635228,0.0014301721,0.19931258,0.32209656,0.04320621],"genre_scores_gemma":[0.33963603,0.00063470664,0.2633635,0.0008228171,0.00043110497,0.00086832,0.3581426,0.008194186,0.027906751],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9962288,0.0013509428,0.00026734654,0.0009036143,0.0009649303,0.00028443124],"domain_scores_gemma":[0.99546415,0.0014145474,0.0001456428,0.0010116239,0.0015049652,0.00045909884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049932385,0.0012122691,0.0013054743,0.003017774,0.0015265471,0.0021457768,0.002348395,0.0017125027,0.016080588],"category_scores_gemma":[0.014696678,0.0006953156,0.00049963075,0.001617188,0.00056102645,0.004493872,0.0033766287,0.0015354735,0.016687606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025257159,0.0003242344,0.0064548464,0.0011200472,0.00029793542,0.00064873265,0.0015239018,0.006115405,0.021274114,0.0053372937,0.72585195,0.22852574],"study_design_scores_gemma":[0.0010469995,0.0006743804,0.017347451,0.00029043158,0.00028927572,0.0008512663,0.001388609,0.22498743,0.09266529,0.018824823,0.6412917,0.00034229248],"about_ca_topic_score_codex":0.018572712,"about_ca_topic_score_gemma":0.020208757,"teacher_disagreement_score":0.018572712,"about_ca_system_score_codex":0.0014624372,"about_ca_system_score_gemma":0.0022705432,"threshold_uncertainty_score":0.05379492},"labels":[],"label_agreement":null},{"id":"W2806219291","doi":"","title":"WIP Event Detection System in TAC KBP 2015 Event Nugget Track.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Event (particle physics); Computer science; Real-time computing; Operating system; Physics","score_opus":0.013471979445909903,"score_gpt":0.25960173772830514,"score_spread":0.24612975828239522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806219291","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07707967,0.002396525,0.27887636,0.0016195545,0.0011918835,0.0015255163,0.22944124,0.37136874,0.036500487],"genre_scores_gemma":[0.2872953,0.0007674349,0.19148937,0.00063837616,0.0002648544,0.0013416301,0.4856939,0.0057880674,0.026721079],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859995,0.0001638987,0.00014841041,0.00042281218,0.0005533734,0.00011144479],"domain_scores_gemma":[0.9971071,0.0006596788,0.00021551337,0.00050398597,0.0012915914,0.00022213788],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017731113,0.0010399405,0.001114642,0.0038056185,0.0009881228,0.001801623,0.0018706226,0.0010152139,0.012156672],"category_scores_gemma":[0.0077303867,0.00049754867,0.00039499506,0.0022275958,0.00021613603,0.0037217403,0.0015844075,0.0012308764,0.012485936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016921392,0.00028517007,0.011790323,0.0013735422,0.00020817597,0.0005966136,0.0007085931,0.0056981035,0.015646577,0.0033158371,0.6680231,0.29066184],"study_design_scores_gemma":[0.000520067,0.00059012236,0.02998163,0.00029662065,0.00036129475,0.000711302,0.0013068884,0.34856477,0.06410485,0.014115722,0.53918904,0.00025770435],"about_ca_topic_score_codex":0.014852383,"about_ca_topic_score_gemma":0.015197161,"teacher_disagreement_score":0.014852383,"about_ca_system_score_codex":0.00075366406,"about_ca_system_score_gemma":0.0015967817,"threshold_uncertainty_score":0.04066807},"labels":[],"label_agreement":null},{"id":"W2806238373","doi":"","title":"Overview of TAC KBP 2015 Event Nugget Track.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Computer science; Event (particle physics); Physics; Operating system","score_opus":0.03524716365777896,"score_gpt":0.30696015754248296,"score_spread":0.271712993884704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806238373","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013045758,0.015241055,0.27220583,0.011022281,0.0047097825,0.004196831,0.3464201,0.17822194,0.15493642],"genre_scores_gemma":[0.028807191,0.004579799,0.14098848,0.0017132972,0.000831396,0.0025014654,0.7524266,0.013322771,0.054829005],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9942154,0.0010996351,0.00044736336,0.00094122556,0.0029433996,0.00035291532],"domain_scores_gemma":[0.98746526,0.0028996044,0.00049880374,0.0019892773,0.005839779,0.0013074508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010876304,0.0022928754,0.0026041463,0.011765129,0.003218965,0.009216939,0.00603569,0.0023918618,0.06060578],"category_scores_gemma":[0.025584107,0.0016491369,0.0013298044,0.009747802,0.0007339686,0.014874795,0.0057194633,0.0036803128,0.055699658],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000369142,0.00019463347,0.0010804534,0.00093289133,0.00009722197,0.00012892945,0.00032892686,0.0019401561,0.0022544765,0.0077613606,0.8723721,0.112539776],"study_design_scores_gemma":[0.0001678347,0.00016592778,0.0022402045,0.00035504287,0.000107180895,0.00023273824,0.00029284935,0.020468976,0.0035607354,0.011841062,0.9604649,0.000102541984],"about_ca_topic_score_codex":0.04852623,"about_ca_topic_score_gemma":0.04918781,"teacher_disagreement_score":0.06060578,"about_ca_system_score_codex":0.003850638,"about_ca_system_score_gemma":0.0077186255,"threshold_uncertainty_score":0.20274651},"labels":[],"label_agreement":null},{"id":"W2806244498","doi":"","title":"Improving DISCERN with Deep Learning.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence","score_opus":0.005620532110856783,"score_gpt":0.2114801175536341,"score_spread":0.20585958544277733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806244498","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040597323,0.011951805,0.9237406,0.0033045844,0.0006306765,0.000092274015,0.0017458566,0.0068731983,0.011063616],"genre_scores_gemma":[0.6050356,0.0029444054,0.3661878,0.0012609197,0.0008355998,0.0001536711,0.0067859124,0.00076731056,0.016028825],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987413,0.00040718744,0.00007614605,0.00031173314,0.00033257544,0.00013099141],"domain_scores_gemma":[0.99613094,0.002396623,0.00017946861,0.0006948396,0.0004210607,0.00017702152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018016729,0.0009585986,0.0014334394,0.0023613474,0.0005973485,0.0016800353,0.0019593178,0.0014713837,0.0054976395],"category_scores_gemma":[0.011382091,0.0004262186,0.00085597904,0.0023162155,0.0007189949,0.0056698946,0.0024360674,0.0033267778,0.003051043],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004184984,0.00028417172,0.0037527396,0.00039179734,0.0001578914,0.00006444819,0.00021030783,0.054878265,0.002581398,0.05366503,0.06661406,0.8169814],"study_design_scores_gemma":[0.0000718472,0.000098673634,0.0007550094,0.000084713385,0.00007592674,0.00008518667,0.00011406914,0.78724575,0.0037341798,0.18779649,0.019917855,0.000020290534],"about_ca_topic_score_codex":0.004283499,"about_ca_topic_score_gemma":0.009198768,"teacher_disagreement_score":0.0054976395,"about_ca_system_score_codex":0.0010750582,"about_ca_system_score_gemma":0.0013499935,"threshold_uncertainty_score":0.01839143},"labels":[],"label_agreement":null},{"id":"W2806318393","doi":"","title":"Sentences Embedding for Slot Filling via Convolutional Neural Networks.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Embedding; Convolutional neural network; Computer science; Artificial intelligence","score_opus":0.014616102968856641,"score_gpt":0.2577463917833841,"score_spread":0.24313028881452745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806318393","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038969178,0.0019564927,0.93581444,0.000893938,0.00059240626,0.00018934983,0.006179612,0.009627044,0.00577745],"genre_scores_gemma":[0.44866937,0.00095007994,0.5184415,0.00033646237,0.00040492095,0.0003615342,0.01640264,0.0008301639,0.013603324],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99946123,0.00014943142,0.0000469275,0.00016151564,0.00011083987,0.000070057664],"domain_scores_gemma":[0.9989016,0.00054783095,0.00007910217,0.0002199207,0.00019362418,0.000057884743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007183676,0.0006954548,0.0006451058,0.0012120538,0.0004121187,0.0008464076,0.0011835345,0.0008359983,0.008935159],"category_scores_gemma":[0.0039452044,0.0003524564,0.00076555373,0.001304152,0.00037895568,0.003203648,0.0013529307,0.0011129424,0.003360778],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058541424,0.00028235323,0.0013267996,0.00056464534,0.000100448575,0.00016678999,0.00041922898,0.025543973,0.011738691,0.04241799,0.0666724,0.8501814],"study_design_scores_gemma":[0.000057241334,0.00011837979,0.00073168013,0.00010095871,0.00006961239,0.00014083469,0.0002079848,0.8435337,0.010652154,0.11669048,0.027662575,0.000034354907],"about_ca_topic_score_codex":0.0028736934,"about_ca_topic_score_gemma":0.006850334,"teacher_disagreement_score":0.008935159,"about_ca_system_score_codex":0.0005808934,"about_ca_system_score_gemma":0.0010611785,"threshold_uncertainty_score":0.029891074},"labels":[],"label_agreement":null},{"id":"W2806332557","doi":"","title":"Events Detection, Coreference and Sequencing: What's next? Overview of the TAC KBP 2017 Event Track.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Coreference; Track (disk drive); Computer science; Event (particle physics); Artificial intelligence; Resolution (logic); Operating system","score_opus":0.056453023116027075,"score_gpt":0.30196327155157776,"score_spread":0.2455102484355507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806332557","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023590064,0.040968765,0.76587015,0.010958002,0.0030454358,0.0016113435,0.059448697,0.05139479,0.043112844],"genre_scores_gemma":[0.07669095,0.023240695,0.53179264,0.0029557615,0.0015199408,0.0024540594,0.3177264,0.008066282,0.035553344],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99396306,0.00118882,0.0006228457,0.0010018714,0.0028339885,0.0003892615],"domain_scores_gemma":[0.9776372,0.0072176238,0.0011614759,0.00288613,0.010083829,0.0010137503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013594358,0.0015848578,0.0016149277,0.011972807,0.0025074387,0.005608699,0.0034846214,0.0022308168,0.008933191],"category_scores_gemma":[0.024548149,0.0011902064,0.0010553785,0.012180241,0.00097439886,0.00997363,0.00423557,0.003816125,0.009335927],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038318575,0.00023059554,0.005946228,0.0017493643,0.00016452392,0.00020312409,0.0011376374,0.00392292,0.007411066,0.010135495,0.2720789,0.696637],"study_design_scores_gemma":[0.00012189831,0.0002637059,0.010442342,0.0011920367,0.00028199903,0.0008230047,0.0009666756,0.04436619,0.019727211,0.03284436,0.888752,0.00021875482],"about_ca_topic_score_codex":0.03868213,"about_ca_topic_score_gemma":0.049969774,"teacher_disagreement_score":0.03868213,"about_ca_system_score_codex":0.0024325289,"about_ca_system_score_gemma":0.0095647015,"threshold_uncertainty_score":0.07691395},"labels":[],"label_agreement":null},{"id":"W2806375004","doi":"","title":"Modeling Event Extraction via Multilingual Data Sources.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Event data; Extraction (chemistry); Event (particle physics); Data extraction; Natural language processing; Information extraction; Chromatography; Chemistry; Physics; MEDLINE","score_opus":0.04752959636667484,"score_gpt":0.31679675028118726,"score_spread":0.2692671539145124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806375004","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0128683,0.0023057768,0.9714556,0.0009845357,0.00021707594,0.00020576676,0.005731155,0.003000639,0.0032311268],"genre_scores_gemma":[0.3789668,0.0028753034,0.5809298,0.00034414037,0.00041673647,0.0006174851,0.030587707,0.0006971206,0.0045647947],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99606836,0.0017061336,0.0004205065,0.000883419,0.00071031303,0.00021132905],"domain_scores_gemma":[0.988325,0.008721796,0.0006703237,0.0009977663,0.0011092812,0.00017580391],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005090906,0.0012544111,0.0010514485,0.005239279,0.00094495446,0.00442834,0.0019551942,0.0013505555,0.0033090587],"category_scores_gemma":[0.025391728,0.0009498033,0.0024916218,0.0065587573,0.00058869494,0.007841467,0.0031869295,0.0019186018,0.0022329798],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012943554,0.00040614876,0.0265634,0.002497423,0.0016183748,0.0016367468,0.002894827,0.22626485,0.0066672107,0.16158488,0.04908957,0.5194822],"study_design_scores_gemma":[0.000063283114,0.000047144502,0.0026978273,0.00020265511,0.00028441724,0.00036071546,0.00043914348,0.8008928,0.004494092,0.15541013,0.0350503,0.000057514306],"about_ca_topic_score_codex":0.012548282,"about_ca_topic_score_gemma":0.01758039,"teacher_disagreement_score":0.012548282,"about_ca_system_score_codex":0.001502194,"about_ca_system_score_gemma":0.0022138793,"threshold_uncertainty_score":0.026923597},"labels":[],"label_agreement":null},{"id":"W2806399785","doi":"","title":"Event Nugget Detection, Classification and Coreference Resolution using Deep Neural Networks and eXtreme Grandient Boosting.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Coreference; Boosting (machine learning); Computer science; Artificial intelligence; Event (particle physics); Artificial neural network; Resolution (logic); Deep neural networks; Pattern recognition (psychology); Physics","score_opus":0.0550547200623276,"score_gpt":0.26604136383774074,"score_spread":0.21098664377541315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806399785","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04455399,0.0018178668,0.945703,0.000537295,0.0002360809,0.0001132182,0.0005168632,0.002469963,0.004051738],"genre_scores_gemma":[0.6022632,0.0006689662,0.38102847,0.0003933013,0.00044772105,0.00017698896,0.0029286672,0.00028100485,0.011811598],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987919,0.0003892991,0.00007069509,0.00035011067,0.00022977806,0.00016828069],"domain_scores_gemma":[0.9978635,0.0010922058,0.00017381365,0.00039292153,0.000366034,0.00011159031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002647175,0.00079332,0.0013440282,0.0020857041,0.0008798258,0.0013742737,0.0022472388,0.001453712,0.002454634],"category_scores_gemma":[0.0057582366,0.0004966361,0.0009459628,0.0017442037,0.00044489533,0.0020187746,0.0020755394,0.0022231066,0.0017085391],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008229591,0.00039362442,0.004365352,0.00020021254,0.00022917251,0.00031940916,0.00029864773,0.07691407,0.014196281,0.014363028,0.026956197,0.860941],"study_design_scores_gemma":[0.000019353905,0.000040222763,0.0011092768,0.00001702991,0.000041564013,0.00009398198,0.000050450264,0.9716813,0.005557558,0.018532934,0.002842905,0.000013468479],"about_ca_topic_score_codex":0.0033823845,"about_ca_topic_score_gemma":0.0073373336,"teacher_disagreement_score":0.0033823845,"about_ca_system_score_codex":0.00070462015,"about_ca_system_score_gemma":0.00082471827,"threshold_uncertainty_score":0.01399976},"labels":[],"label_agreement":null},{"id":"W2806477841","doi":"10.3390/info9060133","title":"A Machine Learning Filter for the Slot Filling Task","year":2018,"lang":"en","type":"article","venue":"Information","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Polytechnique Montréal","funders":"Fonds de recherche du Québec – Nature et technologies","keywords":"Computer science; Classifier (UML); Artificial intelligence; USable; Relationship extraction; Natural language processing; Precision and recall; Filter (signal processing); Information extraction; Machine learning; Speech recognition; World Wide Web","score_opus":0.02177142452139717,"score_gpt":0.24031188179576599,"score_spread":0.21854045727436883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806477841","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037166946,0.00081936904,0.9426866,0.00055373536,0.00027792918,0.00040857232,0.0017852293,0.014282868,0.0020187688],"genre_scores_gemma":[0.21729966,0.00033641164,0.7666698,0.00046393913,0.00034859622,0.0008130048,0.006293677,0.0004953519,0.0072795493],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997675,0.00039678763,0.00025239476,0.0007995999,0.0006228609,0.00025335336],"domain_scores_gemma":[0.99298483,0.0045434604,0.00028781613,0.0005916577,0.0014561706,0.0001360712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004772696,0.001744526,0.0019507916,0.0038855171,0.0017956644,0.0021686947,0.0020986253,0.0037949856,0.005125241],"category_scores_gemma":[0.0104269255,0.0005510424,0.0016048874,0.0028723096,0.0005495237,0.002705075,0.0009830344,0.0020939321,0.0052943053],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009238926,0.00048628668,0.005734053,0.0004021507,0.00018529981,0.00031331897,0.00028331587,0.016045496,0.029372446,0.0045714984,0.027466059,0.91421616],"study_design_scores_gemma":[0.0001074432,0.00033621292,0.0045892755,0.000075532385,0.00016464265,0.00048180055,0.0001500533,0.91157186,0.047382414,0.007199036,0.027868481,0.00007324111],"about_ca_topic_score_codex":0.010263569,"about_ca_topic_score_gemma":0.010316349,"teacher_disagreement_score":0.010263569,"about_ca_system_score_codex":0.0013313853,"about_ca_system_score_gemma":0.002543913,"threshold_uncertainty_score":0.02524072},"labels":[],"label_agreement":null},{"id":"W2806484018","doi":"","title":"Stanford at TAC KBP 2017: Building a Trilingual Relational Knowledge Graph.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Graph; Knowledge graph; Natural language processing; Artificial intelligence; Theoretical computer science","score_opus":0.02639469434775506,"score_gpt":0.29917357277583595,"score_spread":0.2727788784280809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806484018","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020076763,0.001893335,0.4727309,0.0031822901,0.0007580515,0.0012000215,0.3264538,0.124713466,0.04899133],"genre_scores_gemma":[0.070042156,0.0011895125,0.41830584,0.00050336926,0.00010579174,0.0008525387,0.4892575,0.006069746,0.013673454],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998336,0.00044010556,0.00015565532,0.00046289276,0.0005026724,0.00010269916],"domain_scores_gemma":[0.9964923,0.0013444333,0.00014598537,0.0009040304,0.0008659585,0.00024724175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016827896,0.0013442674,0.000898247,0.007023608,0.0018264549,0.0034569537,0.0025434738,0.0013918842,0.036546256],"category_scores_gemma":[0.013718636,0.0011778201,0.0014498057,0.005142799,0.00068473595,0.0077208495,0.0041154996,0.002847106,0.016558541],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030778372,0.00034411935,0.0025311382,0.0016360657,0.00027217824,0.00062572793,0.0015130528,0.011657095,0.0032706524,0.048703983,0.6978754,0.23126285],"study_design_scores_gemma":[0.00023435394,0.000112487156,0.0023572992,0.00056754996,0.000274815,0.0004280149,0.0015749438,0.11799939,0.0053601633,0.119700804,0.75124377,0.00014642939],"about_ca_topic_score_codex":0.042543467,"about_ca_topic_score_gemma":0.07504329,"teacher_disagreement_score":0.042543467,"about_ca_system_score_codex":0.0013821484,"about_ca_system_score_gemma":0.004319556,"threshold_uncertainty_score":0.12225944},"labels":[],"label_agreement":null},{"id":"W2807013141","doi":"","title":"Distributed Neural Embedding for Event Nugget Extraction.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Extraction (chemistry); Embedding; Computer science; Event (particle physics); Artificial intelligence; Chromatography; Chemistry; Physics","score_opus":0.023670185116703112,"score_gpt":0.30738786557704,"score_spread":0.2837176804603369,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807013141","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0346969,0.003110805,0.94837433,0.0006160536,0.00040574305,0.00016650517,0.003788139,0.004735273,0.00410628],"genre_scores_gemma":[0.6297065,0.0018905357,0.3348914,0.00027766597,0.0005922377,0.00037908455,0.017357387,0.0004739334,0.0144312205],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993561,0.00017288595,0.00005451944,0.00021876142,0.00011167541,0.000086054766],"domain_scores_gemma":[0.9987729,0.0006272029,0.00009666119,0.00025524068,0.00019806332,0.000049823586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009902886,0.00090815686,0.00097161176,0.0025960922,0.0005107426,0.0010612311,0.0013967261,0.0011912044,0.003489982],"category_scores_gemma":[0.004132834,0.0003356708,0.0008030472,0.003103533,0.00034400306,0.002709916,0.0013638163,0.0015570297,0.0025153486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006775597,0.0003421885,0.0035453762,0.0004954242,0.00020315638,0.00022291414,0.00023450976,0.04448191,0.009273128,0.018005267,0.036110524,0.8864081],"study_design_scores_gemma":[0.000027559625,0.00007219676,0.0020368497,0.00006321379,0.0000803056,0.0001521371,0.00013684617,0.93629986,0.0049607595,0.04499243,0.011152116,0.000025727943],"about_ca_topic_score_codex":0.0034904263,"about_ca_topic_score_gemma":0.0070544267,"teacher_disagreement_score":0.0034904263,"about_ca_system_score_codex":0.000552446,"about_ca_system_score_gemma":0.00076795334,"threshold_uncertainty_score":0.011675119},"labels":[],"label_agreement":null},{"id":"W2807100726","doi":"","title":"ISCAS_Sogou at TAC-KBP 2017.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.01852562386134024,"score_gpt":0.2737249982392757,"score_spread":0.2551993743779355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807100726","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007798737,0.007817722,0.17297705,0.042338528,0.031117294,0.0010919614,0.14584066,0.14994298,0.441075],"genre_scores_gemma":[0.04700767,0.0048772553,0.07527944,0.0033900107,0.009225274,0.00090822065,0.23938158,0.041615788,0.5783147],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99575174,0.0013300842,0.00018501443,0.0007686859,0.0015492457,0.00041525156],"domain_scores_gemma":[0.9919104,0.0024583433,0.00022224207,0.0014047752,0.002371665,0.0016324968],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0072614853,0.002230822,0.002678637,0.003504286,0.002816154,0.010805388,0.0037759282,0.0031981496,0.4424572],"category_scores_gemma":[0.017779658,0.00083412766,0.0013701381,0.0036825761,0.0012743614,0.00937165,0.005325775,0.0038529632,0.36603388],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024015347,0.00010951755,0.00014008603,0.00023093176,0.000021043756,0.00015776708,0.00014499118,0.0005248371,0.0009560032,0.009123287,0.9464154,0.0419359],"study_design_scores_gemma":[0.00009986432,0.000040626488,0.00040225923,0.00014049977,0.000014346935,0.000100119585,0.00019365629,0.0049036904,0.0014840795,0.016546406,0.9760356,0.00003891653],"about_ca_topic_score_codex":0.011651666,"about_ca_topic_score_gemma":0.013883839,"teacher_disagreement_score":0.5575428,"about_ca_system_score_codex":0.003221786,"about_ca_system_score_gemma":0.003349701,"threshold_uncertainty_score":0.7952671},"labels":[],"label_agreement":null},{"id":"W2807172527","doi":"","title":"Event Argument Linking and Event Nugget Detection Task: IHMC DISCERN System Report.","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Argument (complex analysis); Computer science; Task (project management); Real-time computing; Systems engineering; Engineering; Medicine; Physics","score_opus":0.011979760791547113,"score_gpt":0.24872184592790852,"score_spread":0.2367420851363614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807172527","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30185005,0.004093652,0.06966957,0.008350719,0.0023325416,0.0060287523,0.40904215,0.1479285,0.050703995],"genre_scores_gemma":[0.24715316,0.0006148607,0.1586391,0.0013247747,0.00062893046,0.0024440123,0.55076176,0.0029318903,0.03550148],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981236,0.00037739618,0.00015355178,0.00062953285,0.00052806287,0.00018777135],"domain_scores_gemma":[0.9912485,0.004408658,0.00030116504,0.0018411063,0.0015054182,0.0006951218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032582625,0.00204318,0.0019821406,0.0034871309,0.0018104954,0.002710036,0.0027647598,0.003423581,0.02743946],"category_scores_gemma":[0.021696245,0.0006894115,0.0009180252,0.0022263327,0.00066143106,0.004174641,0.0040234774,0.0029418266,0.024022374],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019312032,0.0011775255,0.0072837197,0.0009625716,0.00018650554,0.00048602282,0.00051123137,0.0022641167,0.006697048,0.002695059,0.78224415,0.19356088],"study_design_scores_gemma":[0.0060537,0.0019022573,0.05303157,0.0005445912,0.00096226356,0.0021207638,0.0039857873,0.22953272,0.08187655,0.03218189,0.58728087,0.000527115],"about_ca_topic_score_codex":0.028557245,"about_ca_topic_score_gemma":0.044695713,"teacher_disagreement_score":0.028557245,"about_ca_system_score_codex":0.0016103052,"about_ca_system_score_gemma":0.0042309747,"threshold_uncertainty_score":0.09179407},"labels":[],"label_agreement":null},{"id":"W2807195765","doi":"","title":"CMU CS Event TAC-KBP2016 Event Argument Extraction System.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Argument (complex analysis); Event (particle physics); Computer science; Extraction (chemistry); Real-time computing; Chemistry; Chromatography; Physics","score_opus":0.009045819101575923,"score_gpt":0.25502098529326056,"score_spread":0.24597516619168464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807195765","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010361366,0.0007120782,0.17155777,0.0012643844,0.0006836106,0.0010272247,0.43221274,0.32570064,0.056480102],"genre_scores_gemma":[0.043867413,0.00035762784,0.200904,0.00044635957,0.00030553754,0.0011064786,0.71773756,0.012990684,0.022284297],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99853516,0.00039459503,0.00017507643,0.0003499616,0.00045192937,0.0000931664],"domain_scores_gemma":[0.9961391,0.0015203682,0.00026410323,0.00078098965,0.0010434366,0.00025193603],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019897143,0.002046693,0.0010276115,0.005616089,0.0010443949,0.002370404,0.0021168997,0.0013843278,0.06513788],"category_scores_gemma":[0.011579464,0.000679283,0.00087170786,0.0032780578,0.00029648916,0.004330727,0.002579948,0.0015310855,0.054106385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006209576,0.00015959772,0.0012140931,0.0010927598,0.000089461006,0.00030279643,0.00033253542,0.0016935166,0.0043892707,0.0068724863,0.8829272,0.10030538],"study_design_scores_gemma":[0.00046275355,0.0001129435,0.0037371016,0.00024583217,0.00012621262,0.00052094716,0.00038979147,0.07430635,0.01674477,0.017738614,0.8854916,0.00012306888],"about_ca_topic_score_codex":0.007336302,"about_ca_topic_score_gemma":0.010096742,"teacher_disagreement_score":0.06513788,"about_ca_system_score_codex":0.0009875172,"about_ca_system_score_gemma":0.0021555298,"threshold_uncertainty_score":0.2179079},"labels":[],"label_agreement":null},{"id":"W2807207822","doi":"","title":"NAIST Participation in the TAC KBP 2016 Cold Start Slot Filling Task: Combining CNN-based and Bootstrapping-based Methods.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Bootstrapping (finance); Computer science; Task (project management); Artificial intelligence; Pattern recognition (psychology); Mathematics; Econometrics; Engineering","score_opus":0.026476614561299,"score_gpt":0.3072488217754716,"score_spread":0.28077220721417256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807207822","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5286577,0.0109629175,0.13960703,0.012514405,0.008085422,0.0026779748,0.10770128,0.056360167,0.13343315],"genre_scores_gemma":[0.6700463,0.0009629155,0.10019515,0.0020630797,0.0010980401,0.0013950446,0.17264454,0.0027190961,0.04887582],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9940246,0.0023970467,0.0003777278,0.0015606622,0.0010448733,0.00059500796],"domain_scores_gemma":[0.9886303,0.006084364,0.00037666544,0.001816926,0.0022758062,0.0008160595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00607316,0.002238773,0.0017280517,0.0024678232,0.0024188198,0.0031871642,0.0033930808,0.0034366078,0.014295862],"category_scores_gemma":[0.029213078,0.00059228897,0.0008485581,0.0020462312,0.00062163273,0.006417051,0.0038975263,0.003190385,0.017807558],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035349573,0.0011747197,0.012385195,0.0014567306,0.00031707887,0.0008445938,0.0020342683,0.009753865,0.009678371,0.0024413662,0.4397643,0.51661456],"study_design_scores_gemma":[0.0009263249,0.0010279923,0.025850138,0.0008161963,0.0005681378,0.0012963025,0.008069435,0.56180197,0.029940959,0.025417883,0.3438559,0.00042880338],"about_ca_topic_score_codex":0.027309233,"about_ca_topic_score_gemma":0.05963565,"teacher_disagreement_score":0.027309233,"about_ca_system_score_codex":0.0013879165,"about_ca_system_score_gemma":0.0033204248,"threshold_uncertainty_score":0.054300547},"labels":[],"label_agreement":null},{"id":"W2807318936","doi":"","title":"KYOTOU at TAC KBP 2017 Event Track: Neural Network-based Event Sequence Classification Model.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Track (disk drive); Computer science; Sequence (biology); Artificial neural network; Artificial intelligence; Operating system; Genetics; Biology","score_opus":0.0454621363911671,"score_gpt":0.3052338252223662,"score_spread":0.2597716888311991,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807318936","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2941489,0.004055846,0.39290747,0.0046194876,0.0028997227,0.0009317371,0.24471343,0.028739918,0.0269835],"genre_scores_gemma":[0.5228523,0.000993404,0.11525634,0.00034644367,0.00045766687,0.0007169421,0.33155477,0.0005672317,0.027254844],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994961,0.000101458136,0.000030738745,0.0001856462,0.00011815876,0.00006785112],"domain_scores_gemma":[0.9992404,0.00017309868,0.00005167277,0.0001305545,0.0003412316,0.00006299942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00093663664,0.00092954835,0.00064205006,0.0015006806,0.0005718632,0.0008841652,0.0014812964,0.0010293655,0.0053424654],"category_scores_gemma":[0.0032152722,0.00029895664,0.00052801543,0.0018213132,0.00014979411,0.0017083153,0.000786119,0.0015725865,0.004688954],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017615695,0.0010279526,0.02500378,0.00044786202,0.0002927697,0.0002368492,0.00017847496,0.08002579,0.008677848,0.004814573,0.43191427,0.44561833],"study_design_scores_gemma":[0.00007979543,0.0001539061,0.012727509,0.000042261545,0.000093184666,0.00007056038,0.00009156822,0.94781756,0.005107004,0.004508038,0.029264798,0.000043890435],"about_ca_topic_score_codex":0.040511064,"about_ca_topic_score_gemma":0.0501157,"teacher_disagreement_score":0.040511064,"about_ca_system_score_codex":0.00096958334,"about_ca_system_score_gemma":0.0012617854,"threshold_uncertainty_score":0.08055055},"labels":[],"label_agreement":null},{"id":"W2807338375","doi":"","title":"OpenIE for Slot Filling at TAC KBP 2017 - System Description.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.03540506824049892,"score_gpt":0.27534120995808603,"score_spread":0.23993614171758712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807338375","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0028672258,0.0005438245,0.45714444,0.0014529411,0.0011446726,0.00039884154,0.21218471,0.28249195,0.04177133],"genre_scores_gemma":[0.08395022,0.0007867724,0.3561636,0.0012660192,0.00045177672,0.0015596645,0.45616853,0.05883263,0.04082076],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978841,0.00050190586,0.0002613248,0.00040066458,0.00068786775,0.0002641794],"domain_scores_gemma":[0.9950054,0.0021600062,0.00022024728,0.0014971404,0.0009269901,0.00019020395],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026837194,0.001490585,0.0010261344,0.0029895653,0.0009395324,0.004588189,0.0023247877,0.002047283,0.14164135],"category_scores_gemma":[0.015924929,0.0011902414,0.0013928947,0.0022356075,0.0006382327,0.0076603773,0.0043232613,0.0026656305,0.07723085],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000706311,0.00010130534,0.0013575294,0.0016027367,0.000089253255,0.000577544,0.0011400233,0.0054606064,0.0031202,0.09365543,0.7793411,0.11284795],"study_design_scores_gemma":[0.000077683806,0.000039723614,0.0004828797,0.0003714229,0.00003236836,0.00021261962,0.00028692553,0.03676772,0.00606346,0.07565391,0.8799357,0.000075677366],"about_ca_topic_score_codex":0.00683014,"about_ca_topic_score_gemma":0.008908937,"teacher_disagreement_score":0.14164135,"about_ca_system_score_codex":0.0017957439,"about_ca_system_score_gemma":0.0027509702,"threshold_uncertainty_score":0.4738375},"labels":[],"label_agreement":null},{"id":"W2807375040","doi":"","title":"Ict_spring Belief and Sentiment System at TAC 2017.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Spring (device); Computer science; Information and Communications Technology; World Wide Web; Engineering","score_opus":0.012717429895996961,"score_gpt":0.25043003606355,"score_spread":0.23771260616755305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807375040","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36213878,0.0007508033,0.052106734,0.0061011654,0.004100387,0.0014547578,0.3361281,0.12724434,0.10997497],"genre_scores_gemma":[0.51417184,0.00017539911,0.04541271,0.0006188414,0.00060993986,0.0007151545,0.3865513,0.0033854104,0.048359413],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991891,0.0002376458,0.000041489766,0.00017011515,0.00025716747,0.00010457118],"domain_scores_gemma":[0.9976858,0.0004870627,0.000076588934,0.0003475735,0.0010240163,0.0003789542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00245216,0.0007410418,0.0004803746,0.0015714325,0.00093771296,0.001396264,0.001017266,0.00078698917,0.01422466],"category_scores_gemma":[0.0066898274,0.00027619425,0.00044117888,0.0010126358,0.0002498335,0.0020631333,0.0010354423,0.0013787755,0.012172469],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012006989,0.0006127403,0.01846172,0.00020061177,0.00009804778,0.00012695127,0.0005455186,0.0072704423,0.0032881098,0.0021919627,0.8736001,0.09240302],"study_design_scores_gemma":[0.00053138885,0.0005527525,0.06977843,0.00011225265,0.00018703344,0.00013492625,0.001259625,0.46051627,0.017187769,0.013601177,0.4359099,0.00022843308],"about_ca_topic_score_codex":0.026301196,"about_ca_topic_score_gemma":0.0562698,"teacher_disagreement_score":0.026301196,"about_ca_system_score_codex":0.0012727348,"about_ca_system_score_gemma":0.0013375685,"threshold_uncertainty_score":0.05229622},"labels":[],"label_agreement":null},{"id":"W2807410225","doi":"","title":"HITS at TAC KBP 2015: Entity discovery and linking, and event nugget detection","year":2015,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Computer science; Physics","score_opus":0.011266315134734087,"score_gpt":0.2500608777063872,"score_spread":0.2387945625716531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807410225","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19412573,0.0077975946,0.21502532,0.006984843,0.005821949,0.0017119778,0.3935547,0.14417762,0.030800251],"genre_scores_gemma":[0.2292082,0.0009805298,0.19148017,0.00071813207,0.000740123,0.0009933813,0.55289817,0.004361602,0.018619752],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9887806,0.0030511455,0.00094398734,0.0019280256,0.0046179,0.0006783191],"domain_scores_gemma":[0.9855805,0.0060284445,0.0005929455,0.0032260125,0.0035883128,0.0009837816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077399747,0.0029580703,0.0019086592,0.010009138,0.0030655493,0.004777387,0.0035403075,0.003849967,0.011718628],"category_scores_gemma":[0.034760736,0.0011378863,0.0015234795,0.0075637675,0.0009916646,0.0069737807,0.004875704,0.003931986,0.009979452],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017598993,0.00067363825,0.010076307,0.001624005,0.0005446982,0.00118314,0.0014237089,0.016482921,0.0060444092,0.008720176,0.7908676,0.16059946],"study_design_scores_gemma":[0.0012039382,0.0006103755,0.03251888,0.00051021844,0.00074136513,0.0019575008,0.0026834174,0.4449214,0.03370876,0.038834773,0.44182625,0.00048307012],"about_ca_topic_score_codex":0.052812498,"about_ca_topic_score_gemma":0.07906974,"teacher_disagreement_score":0.052812498,"about_ca_system_score_codex":0.002040169,"about_ca_system_score_gemma":0.004069784,"threshold_uncertainty_score":0.10501015},"labels":[],"label_agreement":null},{"id":"W2807512022","doi":"10.63317/2qxky84mur43","title":"PhotoshopQuiA: A Corpus of Non-Factoid Questions and Answers for Why-Question Answering","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Question answering; Computer science; Natural language processing; Information retrieval; Artificial intelligence","score_opus":0.017157954941915375,"score_gpt":0.2741351856809426,"score_spread":0.25697723073902723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807512022","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06520517,0.005879145,0.05341606,0.0029578004,0.00093670364,0.0020636087,0.8053351,0.015580822,0.048625663],"genre_scores_gemma":[0.10110221,0.0013502045,0.06093398,0.0007091564,0.00020701929,0.002349704,0.8163957,0.002128084,0.014823955],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985801,0.0005166791,0.000121396675,0.00037583648,0.00032099974,0.00008489269],"domain_scores_gemma":[0.9934716,0.0043913517,0.00021012203,0.0007090334,0.00094885484,0.000269072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012589738,0.0012420807,0.00061399746,0.0037460916,0.0014437311,0.0012864389,0.0014168,0.0016071581,0.029499702],"category_scores_gemma":[0.012102375,0.00052394194,0.0006947099,0.003375719,0.00081594463,0.0026482453,0.0018195525,0.0019142267,0.011320103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006078966,0.0003081444,0.0034248782,0.0062600467,0.00013974361,0.00076170743,0.0069730286,0.0021294015,0.018885983,0.014844277,0.8294355,0.11622937],"study_design_scores_gemma":[0.00020269693,0.000072060604,0.01602034,0.00046041352,0.000087919616,0.00077907724,0.002729286,0.007328784,0.009226528,0.00902377,0.95396185,0.000107343156],"about_ca_topic_score_codex":0.017641736,"about_ca_topic_score_gemma":0.031929076,"teacher_disagreement_score":0.029499702,"about_ca_system_score_codex":0.0015257552,"about_ca_system_score_gemma":0.002332822,"threshold_uncertainty_score":0.09868634},"labels":[],"label_agreement":null},{"id":"W2807542997","doi":"","title":"IECAS Event Detection System at TAC KBP 2017 Event Nugget Track.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Event (particle physics); Computer science; Real-time computing; Operating system; Physics","score_opus":0.01379581950313419,"score_gpt":0.26518374894433394,"score_spread":0.25138792944119975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807542997","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.077247195,0.0015978743,0.16162334,0.0027259819,0.0023801709,0.0016639312,0.28996694,0.38546586,0.077328704],"genre_scores_gemma":[0.23072408,0.0005112202,0.11694607,0.00068973866,0.0005793804,0.0009754045,0.6013674,0.0076060346,0.04060061],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978569,0.00033570145,0.00015842526,0.00052119984,0.00096688705,0.00016087574],"domain_scores_gemma":[0.9947745,0.0009132733,0.00031390163,0.0007327808,0.0027833043,0.00048220213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030236132,0.0011104695,0.0010895209,0.0042468086,0.0010164406,0.002670522,0.0015351293,0.0010817008,0.017979506],"category_scores_gemma":[0.009377388,0.00040924194,0.00033604348,0.0021206148,0.0002445425,0.0029253685,0.0014703374,0.0015463297,0.019673126],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013280205,0.00031603864,0.013208191,0.0005100554,0.00015941817,0.00034833344,0.0005887642,0.0038267327,0.009528226,0.0028232143,0.8371094,0.13025364],"study_design_scores_gemma":[0.0006682696,0.00049571635,0.025646783,0.00024126197,0.00024035804,0.00034352037,0.0011483551,0.21374235,0.034207035,0.010160463,0.71290463,0.00020122447],"about_ca_topic_score_codex":0.021099724,"about_ca_topic_score_gemma":0.027408917,"teacher_disagreement_score":0.021099724,"about_ca_system_score_codex":0.0010938862,"about_ca_system_score_gemma":0.0019916801,"threshold_uncertainty_score":0.060147405},"labels":[],"label_agreement":null},{"id":"W2807570555","doi":"","title":"Cornell Belief and Sentiment System at TAC 2016.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing","score_opus":0.008785457471002251,"score_gpt":0.21518863046302392,"score_spread":0.20640317299202166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807570555","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18802586,0.0012398936,0.04204134,0.0062631466,0.0024286292,0.0011622565,0.5767321,0.08628851,0.095818266],"genre_scores_gemma":[0.3173656,0.00024353375,0.04495462,0.000560309,0.00041600765,0.0007159774,0.59496576,0.002152624,0.0386256],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988187,0.000379276,0.00006649265,0.00021975317,0.00039427803,0.00012162085],"domain_scores_gemma":[0.9967308,0.00078023574,0.00011357755,0.0005854085,0.0014597527,0.00033032108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031591437,0.0008556116,0.00055656955,0.00208866,0.0009832358,0.0016881082,0.0011492508,0.0010343194,0.0169684],"category_scores_gemma":[0.011949157,0.0003200373,0.00045564523,0.0014106749,0.0002565567,0.0028372367,0.0011594843,0.0015145418,0.013461499],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060487614,0.00035826417,0.010334675,0.00016185666,0.000087236374,0.0000940872,0.0002774387,0.004566176,0.0011525579,0.0022291255,0.9082897,0.07184405],"study_design_scores_gemma":[0.000533456,0.0003346894,0.040789697,0.00014885003,0.00019113086,0.00011555056,0.00058664184,0.42528147,0.009350608,0.014903078,0.50755954,0.00020529534],"about_ca_topic_score_codex":0.062154867,"about_ca_topic_score_gemma":0.1144866,"teacher_disagreement_score":0.062154867,"about_ca_system_score_codex":0.0021208045,"about_ca_system_score_gemma":0.0018062968,"threshold_uncertainty_score":0.12358618},"labels":[],"label_agreement":null},{"id":"W2807572306","doi":"","title":"AIPHES-HD system at TAC KBP 2016: Neural Event Trigger Span Detection and Event Type and Realis Disambiguation with Word Embeddings.","year":2016,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Event (particle physics); Word (group theory); Computer science; Span (engineering); Natural language processing; Artificial intelligence; Linguistics; Physics; Engineering; Philosophy; Astrophysics","score_opus":0.008474831985252954,"score_gpt":0.2336761292935846,"score_spread":0.22520129730833166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807572306","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08737161,0.0023146442,0.13597423,0.0011291449,0.0024781385,0.0009972369,0.25399607,0.49231625,0.023422722],"genre_scores_gemma":[0.2566087,0.0006312348,0.22696678,0.00069738715,0.0005797607,0.0015663783,0.47234926,0.011641768,0.02895867],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99930644,0.000105597916,0.000055461234,0.0003035209,0.00015390494,0.0000750626],"domain_scores_gemma":[0.9988525,0.00026743676,0.000071921306,0.000366205,0.000323547,0.00011844448],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009882055,0.0021279347,0.0012474896,0.0019645034,0.0007596318,0.0016549897,0.0026619544,0.0015495052,0.035509482],"category_scores_gemma":[0.0040612444,0.0007879535,0.00071064394,0.0012534908,0.00035452406,0.004071085,0.002861013,0.0017642085,0.03020438],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026150236,0.0005108956,0.0034655891,0.0011422307,0.00043694096,0.0006628223,0.00040087124,0.008169574,0.026110956,0.0022860013,0.6436742,0.31052494],"study_design_scores_gemma":[0.0024131734,0.0010701378,0.020882266,0.00037107512,0.0006123952,0.0009611169,0.0009556892,0.52286613,0.10259695,0.033586867,0.31315342,0.0005307501],"about_ca_topic_score_codex":0.0150561575,"about_ca_topic_score_gemma":0.019500991,"teacher_disagreement_score":0.035509482,"about_ca_system_score_codex":0.0007271982,"about_ca_system_score_gemma":0.0014100489,"threshold_uncertainty_score":0.1187911},"labels":[],"label_agreement":null},{"id":"W2807583563","doi":"10.48550/arxiv.1806.03489","title":"Robust Lexical Features for Improved Neural Network Named-Entity Recognition","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Named-entity recognition; Word (group theory); Artificial intelligence; Matching (statistics); Representation (politics); Feature vector; Feature (linguistics); Natural language processing; Artificial neural network; Entity linking; Recurrent neural network; Pattern recognition (psychology); Knowledge base; Mathematics","score_opus":0.1321239846210392,"score_gpt":0.200364422793399,"score_spread":0.06824043817235981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807583563","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13278243,0.0048939986,0.8130838,0.0011317735,0.00067815767,0.00020364377,0.008169241,0.032732032,0.0063249418],"genre_scores_gemma":[0.6286769,0.0010143411,0.32393247,0.00031523363,0.00032428346,0.00029447913,0.035773113,0.0008270648,0.008842105],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99868554,0.0003582745,0.00014174322,0.00044365702,0.0002484378,0.00012222654],"domain_scores_gemma":[0.9981026,0.00070707215,0.00014696318,0.00055900676,0.00043220515,0.000052295152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017735484,0.0014628384,0.0011526116,0.002436071,0.00051710237,0.0012824374,0.0020818876,0.0011744329,0.004010779],"category_scores_gemma":[0.0055806767,0.00041917077,0.00083590904,0.002402836,0.00036963666,0.00460742,0.001703997,0.0016794765,0.0043961834],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005856959,0.0004905936,0.0036338815,0.00044241457,0.00021620412,0.00025798555,0.0001291553,0.08316369,0.02514764,0.007649493,0.046555616,0.8317277],"study_design_scores_gemma":[0.00004295548,0.000088389424,0.0012157743,0.000029810146,0.00005988604,0.00007562439,0.00005376573,0.96713656,0.011139482,0.012289342,0.00783738,0.000031050462],"about_ca_topic_score_codex":0.00464859,"about_ca_topic_score_gemma":0.008465516,"teacher_disagreement_score":0.00464859,"about_ca_system_score_codex":0.00067265116,"about_ca_system_score_gemma":0.0006316391,"threshold_uncertainty_score":0.013417363},"labels":[],"label_agreement":null},{"id":"W2809399426","doi":"10.1017/s1366728918000755","title":"Scaling up: How computational models can propel bilingualism research forward","year":2018,"lang":"en","type":"article","venue":"Bilingualism Language and Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Computational model; Word (group theory); Similarity (geometry); Scaling; Empirical research; Artificial intelligence; Natural language processing; Mathematics; Statistics","score_opus":0.08032217639197745,"score_gpt":0.3427641043788983,"score_spread":0.2624419279869209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809399426","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040076133,0.003380159,0.87867606,0.02265397,0.0004794693,0.00013432905,0.00030482505,0.0007012604,0.053593777],"genre_scores_gemma":[0.65348595,0.0025830662,0.3352404,0.0016611504,0.0004905656,0.0005003881,0.00037019353,0.00076968723,0.004898572],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.994262,0.0039004511,0.00022715489,0.00084564893,0.00058094447,0.00018383654],"domain_scores_gemma":[0.9707236,0.021906458,0.0010794089,0.00434765,0.0013636864,0.0005791993],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012386911,0.0017380534,0.0016942394,0.0026553215,0.0016987248,0.007173176,0.0023292855,0.0019566305,0.010412117],"category_scores_gemma":[0.04449186,0.0014280534,0.0023245553,0.0023418409,0.0065503772,0.022131223,0.0053971824,0.0058034454,0.002161963],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004359271,0.000038543512,0.0016345787,0.00015357723,0.000101040016,0.000097478594,0.0021654423,0.017830733,0.0003043829,0.95209444,0.0016985246,0.023837592],"study_design_scores_gemma":[0.000016807324,0.000012999192,0.00017373892,0.000045317887,0.000016755055,0.000029992683,0.00024136886,0.0371736,0.00009198747,0.95779926,0.0043826466,0.000015552774],"about_ca_topic_score_codex":0.0063507813,"about_ca_topic_score_gemma":0.0052370178,"teacher_disagreement_score":0.012386911,"about_ca_system_score_codex":0.0026522535,"about_ca_system_score_gemma":0.0023673985,"threshold_uncertainty_score":0.06550896},"labels":[],"label_agreement":null},{"id":"W2810702571","doi":"10.1609/icwsm.v12i1.15021","title":"Unsupervised Model for Topic Viewpoint Discovery in Online Debates Leveraging Author Interactions","year":2018,"lang":"en","type":"article","venue":"Proceedings of the International AAAI Conference on Web and Social Media","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Viewpoints; Identification (biology); Computer science; Topic model; Cluster analysis; Artificial intelligence; Homophily; Data science; Context (archaeology); Machine learning; Sociology; Social science","score_opus":0.09079919561186636,"score_gpt":0.31117625223726464,"score_spread":0.22037705662539828,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2810702571","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24743327,0.0010551387,0.7452011,0.00048116336,0.00008337398,0.00019948052,0.00067328947,0.0007626845,0.0041104937],"genre_scores_gemma":[0.8504543,0.00038089798,0.14207743,0.000111982044,0.00019629802,0.00030752612,0.0025174718,0.0001347019,0.0038193322],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99850225,0.0005621068,0.00007027216,0.0005584757,0.00019218567,0.000114737646],"domain_scores_gemma":[0.99659735,0.0024012462,0.0002955981,0.0002637746,0.00033365923,0.00010847247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021450375,0.00085831893,0.0009480074,0.0022636694,0.00071076927,0.0016825661,0.0017682987,0.0013907837,0.0012117132],"category_scores_gemma":[0.006326505,0.00049663254,0.0014070483,0.0014995314,0.0007271189,0.0020640464,0.001140714,0.0014958269,0.0009279175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002425184,0.00072878407,0.06908063,0.00070431345,0.00058103347,0.0012586508,0.0070154318,0.35533574,0.039242778,0.052632727,0.008307881,0.46268687],"study_design_scores_gemma":[0.000029803128,0.000040231134,0.0033301485,0.000015644617,0.000046241235,0.00008529563,0.0001569985,0.98405,0.001885466,0.008948906,0.0013952011,0.000016002245],"about_ca_topic_score_codex":0.0038854987,"about_ca_topic_score_gemma":0.0070398357,"teacher_disagreement_score":0.0038854987,"about_ca_system_score_codex":0.0009168648,"about_ca_system_score_gemma":0.00094628375,"threshold_uncertainty_score":0.011344194},"labels":[],"label_agreement":null},{"id":"W2811039307","doi":"10.1007/s13278-018-0523-0","title":"Entity linking of tweets based on dominant entity candidates","year":2018,"lang":"en","type":"article","venue":"Social Network Analysis and Mining","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Thomson Reuters (Canada)","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Entity linking; Computer science; Information retrieval; Annotation; Context (archaeology); Limiting; Process (computing); Space (punctuation); Natural language processing; Named entity; Artificial intelligence; Knowledge base","score_opus":0.013069361562876209,"score_gpt":0.2553194528917092,"score_spread":0.242250091328833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2811039307","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61696875,0.0020903193,0.3407709,0.0009929196,0.0005579721,0.00096771785,0.016785324,0.0020976022,0.018768514],"genre_scores_gemma":[0.8762843,0.00069363724,0.09607716,0.00007109676,0.00032198086,0.00042792273,0.019250887,0.00015085716,0.006722283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99803394,0.00040144005,0.00018334675,0.0005929607,0.0005732989,0.0002150322],"domain_scores_gemma":[0.994354,0.0031900143,0.0005565335,0.0005073384,0.0011411955,0.000251002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014715651,0.00089929905,0.0006573923,0.009274587,0.0018111108,0.0021362891,0.0009541316,0.0011461027,0.0030571118],"category_scores_gemma":[0.008659597,0.00037716175,0.0011030881,0.0070804637,0.00035186906,0.003517602,0.0013087852,0.0010576955,0.0019324642],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029990629,0.0013375615,0.2979504,0.0015541783,0.0012503725,0.0038315037,0.0042987973,0.041207064,0.07336357,0.03486718,0.0381533,0.499187],"study_design_scores_gemma":[0.000075765165,0.0004205992,0.10860348,0.00025088934,0.0013032606,0.002050511,0.0023600531,0.7638738,0.05187451,0.02721097,0.041804757,0.00017137843],"about_ca_topic_score_codex":0.0025845768,"about_ca_topic_score_gemma":0.0056537525,"teacher_disagreement_score":0.009274587,"about_ca_system_score_codex":0.00048096105,"about_ca_system_score_gemma":0.00090818375,"threshold_uncertainty_score":0.010227084},"labels":[],"label_agreement":null},{"id":"W2816262648","doi":"","title":"Farewell Freebase: Migrating the SimpleQuestions Dataset to DBpedia.","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Leverage (statistics); Question answering; Knowledge graph; Information retrieval; Benchmark (surveying); Entity linking; Task (project management); Simple (philosophy); Graph; World Wide Web; Knowledge base; Theoretical computer science; Artificial intelligence","score_opus":0.07890198444019045,"score_gpt":0.36377097578953593,"score_spread":0.2848689913493455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2816262648","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017401932,0.0016951718,0.009498534,0.0011757609,0.00047769304,0.0004464388,0.93909186,0.017922837,0.012289765],"genre_scores_gemma":[0.014149074,0.00028060342,0.01578847,0.00035260763,0.000034194796,0.00027387333,0.9673031,0.00057177624,0.0012462791],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965287,0.00076141173,0.00043797898,0.0011170573,0.00093761674,0.00021724292],"domain_scores_gemma":[0.99260265,0.0023566731,0.00050498825,0.0021675685,0.001708113,0.0006599839],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024799365,0.0020080474,0.0009878329,0.008066264,0.0023137277,0.003215827,0.0036791137,0.0026229224,0.009437288],"category_scores_gemma":[0.01734786,0.0007483185,0.0014942375,0.007305928,0.0009610266,0.005626526,0.003969439,0.0029743367,0.008098693],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034235689,0.00038843465,0.0051211677,0.0031214429,0.00028425342,0.000446297,0.0007706146,0.0039012586,0.0029440888,0.007295785,0.9403528,0.035031643],"study_design_scores_gemma":[0.0003500033,0.000118920114,0.012370725,0.00047614003,0.00012471685,0.00056099176,0.0013848193,0.021178383,0.007407871,0.014216414,0.9416279,0.00018310717],"about_ca_topic_score_codex":0.069076136,"about_ca_topic_score_gemma":0.09977357,"teacher_disagreement_score":0.069076136,"about_ca_system_score_codex":0.002605333,"about_ca_system_score_gemma":0.0032135013,"threshold_uncertainty_score":0.13734812},"labels":[],"label_agreement":null},{"id":"W2847160827","doi":"","title":"Reproducing and Regularizing the SCRN Model","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Dropout (neural networks); Tying; Computer science; Language model; Task (project management); Artificial intelligence; Data modeling; Machine learning; Algorithm; Engineering","score_opus":0.07632542245516838,"score_gpt":0.3238737291177107,"score_spread":0.2475483066625423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2847160827","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045707133,0.00011119291,0.9482759,0.00031986428,0.0001293001,0.000044462508,0.00032015878,0.001803589,0.003288454],"genre_scores_gemma":[0.74690574,0.00019879869,0.23963735,0.00030223955,0.00010086029,0.000190713,0.0014172589,0.0007113001,0.010535665],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995832,0.000087578934,0.000017431232,0.00016138722,0.0000979853,0.000052507206],"domain_scores_gemma":[0.9993351,0.00019739204,0.000061241524,0.00023535547,0.0001356305,0.00003521752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008163284,0.00065150874,0.00048780505,0.00031704782,0.0002457387,0.00045949512,0.0014371332,0.00079189474,0.0028185009],"category_scores_gemma":[0.0035232746,0.00033077117,0.00069771917,0.00033341543,0.0005455735,0.0012201883,0.0010311921,0.0014380923,0.0014877109],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011165804,0.00009218569,0.0014485373,0.00010195495,0.000071367904,0.00028808953,0.00016034198,0.82941943,0.0340237,0.044679325,0.004983946,0.08461945],"study_design_scores_gemma":[0.0000037091102,0.000014216019,0.00009772358,0.000002505546,0.00000576455,0.000027926328,0.000004114226,0.9925249,0.0017200515,0.0050422954,0.00055130036,0.0000054722977],"about_ca_topic_score_codex":0.0057529435,"about_ca_topic_score_gemma":0.0082282415,"teacher_disagreement_score":0.0057529435,"about_ca_system_score_codex":0.0004698087,"about_ca_system_score_gemma":0.0010029012,"threshold_uncertainty_score":0.011438906},"labels":[],"label_agreement":null},{"id":"W2848493808","doi":"10.48550/arxiv.1808.06167","title":"Source-Critical Reinforcement Learning for Transferring Spoken Language Understanding to a New Language","year":2018,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Reinforcement learning; Sentence; Natural language processing; Artificial intelligence; Domain (mathematical analysis); Adaptation (eye); Spoken language; Task (project management); Language model","score_opus":0.1008120889274694,"score_gpt":0.22552977491644294,"score_spread":0.12471768598897354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2848493808","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05879938,0.00015545075,0.9357862,0.00022879109,0.00005214805,0.00009857719,0.000050469735,0.0030951598,0.0017338595],"genre_scores_gemma":[0.81523085,0.00009474191,0.18134502,0.00020109609,0.000042251602,0.0002414509,0.00021044766,0.00024539683,0.0023886845],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992704,0.00033117062,0.00003336326,0.00020295897,0.00010666981,0.000055487115],"domain_scores_gemma":[0.99805284,0.0012434222,0.00012716066,0.00019683743,0.0002916046,0.00008806859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016558598,0.00081685017,0.00069919956,0.00037387456,0.00036611978,0.0005245701,0.0010612296,0.00066053966,0.0021192862],"category_scores_gemma":[0.005565022,0.00032725872,0.00042791793,0.0003206212,0.00082439417,0.0013626951,0.0011436424,0.0014032685,0.0006274018],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038709247,0.00047376793,0.002186503,0.00020898807,0.00008957299,0.0002571102,0.0007072268,0.5146832,0.030393742,0.0068106963,0.003526299,0.44027582],"study_design_scores_gemma":[0.000016274973,0.00004311302,0.00012022218,0.0000037469333,0.000007265948,0.000015062589,0.000021943124,0.993228,0.0036197668,0.002518789,0.00039828106,0.000007500899],"about_ca_topic_score_codex":0.003876319,"about_ca_topic_score_gemma":0.0030815613,"teacher_disagreement_score":0.003876319,"about_ca_system_score_codex":0.00084202754,"about_ca_system_score_gemma":0.0010466225,"threshold_uncertainty_score":0.008757114},"labels":[],"label_agreement":null},{"id":"W2865911429","doi":"","title":"Abstractive Unsupervised Multi-Document Summarization using Paraphrastic Sentence Fusion","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Natural language processing; Artificial intelligence; Set (abstract data type); Machine translation; Word embedding; Multi-document summarization; Word (group theory); Information retrieval; Embedding; Linguistics","score_opus":0.09255504555690312,"score_gpt":0.35067864217008043,"score_spread":0.2581235966131773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2865911429","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018355064,0.0009101945,0.9712451,0.00022128792,0.00011178907,0.00016683199,0.0006225855,0.0072473385,0.0011198493],"genre_scores_gemma":[0.1897727,0.00072458456,0.7970082,0.00034314755,0.000349434,0.0003072501,0.0060743783,0.0005512945,0.0048690033],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988194,0.00032943042,0.0001243749,0.0003543253,0.00031309293,0.00005918457],"domain_scores_gemma":[0.9974045,0.0008127111,0.00038971123,0.00051168137,0.0008006206,0.00008063941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012847204,0.0014694282,0.0013671773,0.0019833946,0.00045852744,0.001104998,0.0013231006,0.0009671329,0.0022918382],"category_scores_gemma":[0.0039166613,0.00035531464,0.0012414479,0.0016408789,0.0003558526,0.0023758374,0.0010276346,0.0012959691,0.002481435],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003737762,0.00033032396,0.0010948894,0.00082284975,0.00029266952,0.00033651828,0.00051307864,0.028880255,0.15129773,0.004211293,0.012095587,0.799751],"study_design_scores_gemma":[0.00010101093,0.0009214451,0.0035138917,0.000079606085,0.00044681405,0.0007933794,0.00033161754,0.7692926,0.18620044,0.012316371,0.02587905,0.00012371186],"about_ca_topic_score_codex":0.0010600868,"about_ca_topic_score_gemma":0.0018231857,"teacher_disagreement_score":0.0022918382,"about_ca_system_score_codex":0.00045356475,"about_ca_system_score_gemma":0.0006494264,"threshold_uncertainty_score":0.0076669455},"labels":[],"label_agreement":null},{"id":"W2875408189","doi":"","title":"The APVA-TURBO Approach To Question Answering in Knowledge Base","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Question answering; Computer science; Correctness; Bottleneck; Knowledge base; Object (grammar); Artificial intelligence; Base (topology); Subject (documents); Information retrieval; Theoretical computer science; Machine learning; Programming language; World Wide Web","score_opus":0.06531592434647242,"score_gpt":0.3406081913215932,"score_spread":0.27529226697512077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2875408189","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004315417,0.00025982465,0.9919853,0.00037404205,0.000028031613,0.000058628513,0.000080165024,0.0019354278,0.0009630673],"genre_scores_gemma":[0.24066839,0.00043367874,0.75274307,0.00058207277,0.00016991462,0.00035281258,0.00080476876,0.00046322547,0.0037820507],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99528235,0.002525405,0.00019102154,0.00091667,0.00082866143,0.00025589118],"domain_scores_gemma":[0.9850085,0.009581022,0.0002882598,0.0034634268,0.001357321,0.00030147578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061624935,0.0010260835,0.0013692237,0.002046456,0.001024182,0.002548789,0.0054737795,0.0027373547,0.0044509736],"category_scores_gemma":[0.022215897,0.0011806963,0.0018609248,0.0020292578,0.0023854258,0.007513308,0.004650464,0.004600807,0.001947391],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004778168,0.00027462177,0.002046547,0.00065059104,0.00023707867,0.00025396133,0.001082996,0.21938737,0.007949492,0.10112928,0.0116391685,0.65487105],"study_design_scores_gemma":[0.000011976727,0.00004874466,0.00017048721,0.000020942216,0.00002477346,0.00008155451,0.000050141414,0.915821,0.0021593964,0.07884306,0.0027545409,0.000013465886],"about_ca_topic_score_codex":0.0065900967,"about_ca_topic_score_gemma":0.007958282,"teacher_disagreement_score":0.0065900967,"about_ca_system_score_codex":0.0014647655,"about_ca_system_score_gemma":0.002045866,"threshold_uncertainty_score":0.032590747},"labels":[],"label_agreement":null},{"id":"W2884872018","doi":"10.1007/s42113-018-0008-2","title":"An Instance Theory of Semantic Memory","year":2018,"lang":"en","type":"article","venue":"Computational Brain & Behavior","topic":"Topic Modeling","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Construct (python library); TRACE (psycholinguistics); Word (group theory); Meaning (existential); Natural language processing; Representation (politics); Artificial intelligence; Context (archaeology); Artifact (error); Natural language; Semantics (computer science); Episodic memory; Language model; Linguistics; Psychology; Cognition","score_opus":0.03040468465680373,"score_gpt":0.30101177783071614,"score_spread":0.2706070931739124,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884872018","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053641327,0.0013370719,0.89834046,0.0070470464,0.00028189312,0.00006238099,0.0007968837,0.00062639214,0.03786649],"genre_scores_gemma":[0.86974007,0.0010816169,0.11463233,0.0010334768,0.0008057649,0.00019256517,0.0012587568,0.00022590856,0.011029525],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9988427,0.00049288967,0.00007836028,0.00028931256,0.0001884443,0.00010820155],"domain_scores_gemma":[0.9963356,0.0022550758,0.00018455443,0.000690705,0.00035854892,0.00017541616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018146434,0.0005869452,0.0010295876,0.001482397,0.0010731826,0.0040919427,0.0021284039,0.001821515,0.009063414],"category_scores_gemma":[0.00973867,0.0005484837,0.001859731,0.0015790203,0.0021227833,0.013458903,0.0016116357,0.0026803508,0.001012986],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030646122,0.00001718505,0.0003912291,0.000039863946,0.00003247153,0.000035231515,0.00021689935,0.0018759818,0.0002730915,0.9854951,0.0016557014,0.009936649],"study_design_scores_gemma":[0.000008247428,0.000008122908,0.00015957127,0.000008440973,0.000016002408,0.00005120698,0.000039689832,0.014041576,0.00013355036,0.9840958,0.0014319359,0.000005872345],"about_ca_topic_score_codex":0.0015537329,"about_ca_topic_score_gemma":0.0010414504,"teacher_disagreement_score":0.009063414,"about_ca_system_score_codex":0.0010364555,"about_ca_system_score_gemma":0.0007672266,"threshold_uncertainty_score":0.030320168},"labels":[],"label_agreement":null},{"id":"W2884915013","doi":"10.1109/icmla.2018.00104","title":"Improving Neural Sequence Labelling Using Additional Linguistic Information","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Chunking (psychology); Natural language processing; Sequence labeling; Artificial intelligence; Sequence (biology); Named-entity recognition; Word (group theory); Labelling; Sentence; Benchmark (surveying); Task (project management); Linguistics","score_opus":0.045526245676960814,"score_gpt":0.26970982756214873,"score_spread":0.22418358188518792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2884915013","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03752179,0.00087636465,0.950531,0.00048701392,0.00022499396,0.00009723633,0.00047105356,0.0063324347,0.003458162],"genre_scores_gemma":[0.4797849,0.00084976753,0.49603888,0.0009054488,0.00022896977,0.00038116527,0.005781451,0.0007609781,0.01526846],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936527,0.00016457115,0.00004262266,0.0002549341,0.00011275968,0.00005994255],"domain_scores_gemma":[0.9969193,0.0015210402,0.00019102468,0.0005875112,0.0006736013,0.00010766177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013706661,0.0012144257,0.0011242148,0.001196493,0.00063657085,0.0010133991,0.0021004134,0.0017887547,0.0032352125],"category_scores_gemma":[0.0062861894,0.00052290224,0.0009987202,0.0012227695,0.000628621,0.004516974,0.0012931513,0.0023815813,0.0030998383],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032274742,0.0003984226,0.0024815607,0.00033010053,0.00010442884,0.00020008552,0.0003695879,0.2625713,0.020880677,0.011713972,0.015409192,0.6852179],"study_design_scores_gemma":[0.000011364447,0.000045434383,0.0002524266,0.000021874712,0.000022862889,0.00004669717,0.000032826243,0.98137075,0.0042852936,0.011471197,0.0024249414,0.000014296727],"about_ca_topic_score_codex":0.007427185,"about_ca_topic_score_gemma":0.015318425,"teacher_disagreement_score":0.007427185,"about_ca_system_score_codex":0.0011555314,"about_ca_system_score_gemma":0.0016305326,"threshold_uncertainty_score":0.014767885},"labels":[],"label_agreement":null},{"id":"W2885080048","doi":"10.1109/ipdpsw.2018.00047","title":"GraphNER: Using Corpus Level Similarities and Graph Propagation for Named Entity Recognition","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; Simon Fraser University","funders":"","keywords":"Conditional random field; Named-entity recognition; Computer science; Natural language processing; Artificial intelligence; Graph; Task (project management); Sequence labeling; Information retrieval; Machine learning; Theoretical computer science","score_opus":0.13108791823064653,"score_gpt":0.2852154965909825,"score_spread":0.15412757836033597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885080048","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012676127,0.00060336594,0.9348133,0.0004552862,0.00023010232,0.0003223099,0.0052020364,0.04345628,0.0022411009],"genre_scores_gemma":[0.120499544,0.00066235673,0.84652525,0.0003708445,0.00015871075,0.0005251597,0.022306051,0.0035806512,0.005371392],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972807,0.00069435936,0.0001728014,0.0011474573,0.000585635,0.000118988115],"domain_scores_gemma":[0.99241567,0.0039533037,0.0007833777,0.0017612554,0.0008961106,0.00019020669],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003432151,0.0018645648,0.0012317906,0.007822518,0.0013636838,0.0021691322,0.0028468894,0.0026205047,0.005207073],"category_scores_gemma":[0.014488334,0.0013224323,0.0020350388,0.006223943,0.0012364361,0.0089414315,0.0030019174,0.0023936133,0.004439222],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005763067,0.00042361495,0.004858136,0.0009925793,0.00053957105,0.00076463365,0.000729012,0.14221267,0.010547246,0.040816117,0.06025355,0.73728657],"study_design_scores_gemma":[0.000068931324,0.00007820775,0.0009783532,0.00005462624,0.00006709409,0.00027490925,0.00011278284,0.90978616,0.007925535,0.060264185,0.020313932,0.00007515709],"about_ca_topic_score_codex":0.013855567,"about_ca_topic_score_gemma":0.025835726,"teacher_disagreement_score":0.013855567,"about_ca_system_score_codex":0.0013898106,"about_ca_system_score_gemma":0.002126443,"threshold_uncertainty_score":0.027549863},"labels":[],"label_agreement":null},{"id":"W2885315449","doi":"10.1111/cogs.12662","title":"Simple Co‐Occurrence Statistics Reproducibly Predict Association Ratings","year":2018,"lang":"en","type":"article","venue":"Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Deutsche Forschungsgemeinschaft","keywords":"Valence (chemistry); Word2vec; Psychology; Statistics; Co-occurrence; Natural language processing; Mathematics; Pattern recognition (psychology); Cognitive psychology; Computer science; Artificial intelligence","score_opus":0.03476346330130259,"score_gpt":0.3255728526597379,"score_spread":0.2908093893584353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885315449","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9505143,0.00048823562,0.043735985,0.00012652406,0.000053764776,0.00007862863,0.0006325039,0.00031394185,0.004056157],"genre_scores_gemma":[0.99075794,0.00008664599,0.00825197,0.000019263289,0.000027387792,0.00004663114,0.00048419752,0.00007072676,0.00025518774],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99163,0.0038540727,0.000869979,0.001711574,0.0017086775,0.0002256402],"domain_scores_gemma":[0.81467974,0.15353559,0.013110329,0.011342529,0.005423537,0.0019082776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010483396,0.0005816516,0.0006926259,0.0019277866,0.00039229603,0.0021481928,0.00044944635,0.0008187964,0.0020375727],"category_scores_gemma":[0.089778155,0.0003176296,0.0006324463,0.0020060157,0.00084350794,0.00291875,0.0009875614,0.00076273974,0.00075317046],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015127085,0.00024066176,0.85479087,0.00062008086,0.0012100606,0.0002563532,0.0026334492,0.006755525,0.030162219,0.0026358515,0.0018180818,0.09736421],"study_design_scores_gemma":[0.0000632958,0.0009882479,0.90931964,0.00005909717,0.0004199577,0.0009228798,0.0012568723,0.06445746,0.011208404,0.008738725,0.0023998322,0.0001656399],"about_ca_topic_score_codex":0.0006026921,"about_ca_topic_score_gemma":0.0010488589,"teacher_disagreement_score":0.010483396,"about_ca_system_score_codex":0.00019652395,"about_ca_system_score_gemma":0.00026121197,"threshold_uncertainty_score":0.055442154},"labels":[],"label_agreement":null},{"id":"W2885406936","doi":"10.48550/arxiv.1812.05168","title":"Searching for Relevant Lessons Learned Using Hybrid Information Retrieval Classifiers: A Case Study in Software Engineering","year":2018,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Information retrieval; Artificial intelligence; Software engineering; Software; Machine learning; Data mining; Data science; Programming language","score_opus":0.17777725301382746,"score_gpt":0.25515918391022635,"score_spread":0.07738193089639889,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885406936","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.973695,0.0016206895,0.020028658,0.0008719104,0.000053586857,0.0003491534,0.0009700107,0.00064661365,0.0017644506],"genre_scores_gemma":[0.93256223,0.0006595607,0.06251349,0.00024004425,0.00007528445,0.00017585543,0.0024728128,0.000057167526,0.0012434938],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994565,0.0025994582,0.0005552304,0.0007739406,0.0012019422,0.0003044004],"domain_scores_gemma":[0.97595066,0.019115416,0.0007251749,0.0014654915,0.0022064466,0.00053681224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007004718,0.00086249976,0.0011702256,0.0038528098,0.0014217299,0.001519088,0.0017812009,0.0021874104,0.0007378855],"category_scores_gemma":[0.020319518,0.0002788701,0.00083111157,0.0049380646,0.00073978974,0.0034806577,0.0011402745,0.0014136059,0.000617278],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023087952,0.006288705,0.17146124,0.0025735206,0.0005962573,0.0042052236,0.006915646,0.05298616,0.015077915,0.0026501848,0.023036607,0.7118997],"study_design_scores_gemma":[0.0006918753,0.005074139,0.08730889,0.0002904212,0.00085341866,0.005409822,0.011417979,0.8032798,0.048910633,0.0060677016,0.030416105,0.00027913318],"about_ca_topic_score_codex":0.009326991,"about_ca_topic_score_gemma":0.013502817,"teacher_disagreement_score":0.009326991,"about_ca_system_score_codex":0.0012256294,"about_ca_system_score_gemma":0.0010284582,"threshold_uncertainty_score":0.037044942},"labels":[],"label_agreement":null},{"id":"W2885533082","doi":"","title":"CL-SciSumm Shared Task - Team Magma.","year":2018,"lang":"en","type":"article","venue":"International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Task (project management); Computer science; Magma; Geology; Engineering; Systems engineering; Volcano; Seismology","score_opus":0.10378663782672867,"score_gpt":0.359354122566273,"score_spread":0.2555674847395444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885533082","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020277066,0.0025091586,0.30045792,0.0041824277,0.009364802,0.0055458224,0.1680821,0.4334396,0.056141198],"genre_scores_gemma":[0.11875079,0.0007060359,0.4038611,0.0024417453,0.0021996067,0.013587362,0.3256307,0.037835572,0.09498705],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9927375,0.0031061142,0.0004233997,0.001853262,0.0010229582,0.0008567456],"domain_scores_gemma":[0.98628044,0.0018591081,0.00031592287,0.005933983,0.0035449823,0.0020656334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010156683,0.0043664644,0.003897162,0.003092716,0.0031220587,0.0055049523,0.008267183,0.004455155,0.13574585],"category_scores_gemma":[0.03876713,0.0018301203,0.0031273821,0.0030719945,0.0008713283,0.005605989,0.010288874,0.0038057577,0.18963212],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030480775,0.00047732954,0.00059916347,0.00094925327,0.00040301195,0.000088362285,0.00023566735,0.0053863525,0.0039651264,0.0024816229,0.8721866,0.11017936],"study_design_scores_gemma":[0.0046822,0.0016163484,0.0019828188,0.0002793195,0.00060755166,0.00023762065,0.0005859025,0.18907073,0.02071598,0.05146581,0.7283273,0.0004284055],"about_ca_topic_score_codex":0.011370271,"about_ca_topic_score_gemma":0.02145815,"teacher_disagreement_score":0.13574585,"about_ca_system_score_codex":0.0021553289,"about_ca_system_score_gemma":0.009372372,"threshold_uncertainty_score":0.4541151},"labels":[],"label_agreement":null},{"id":"W2885765530","doi":"10.18653/v1/p19-1041","title":"Disentangled Representation Learning for Non-Parallel Text Style Transfer","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Nvidia","keywords":"Fluency; Representation (politics); Computer science; Style (visual arts); Task (project management); Simple (philosophy); Artificial intelligence; Natural language processing; Space (punctuation); Adversarial system; Latent variable; Transfer (computing); Transfer of learning; Machine learning; Linguistics","score_opus":0.03807439985645557,"score_gpt":0.29549499559505127,"score_spread":0.2574205957385957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885765530","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01864759,0.0003150378,0.97851026,0.00019043173,0.00006418946,0.00003699148,0.000102065474,0.0009168569,0.0012166612],"genre_scores_gemma":[0.7079774,0.0007342121,0.27586445,0.00036639342,0.000422811,0.00034098802,0.0012847552,0.0005050875,0.012503807],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987984,0.0005857987,0.000044317003,0.00030328066,0.00018090263,0.00008738582],"domain_scores_gemma":[0.9974306,0.0013357701,0.00020840798,0.0007291086,0.00018263646,0.000113368034],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023163192,0.0015681761,0.001099959,0.000784562,0.0005056433,0.0011694434,0.0014782569,0.0013449346,0.0029470718],"category_scores_gemma":[0.006244321,0.0006224651,0.0012238905,0.001045226,0.0011804866,0.0036493489,0.0030117247,0.0032552208,0.00181407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039373583,0.0004094712,0.0016022353,0.00021469996,0.00024050516,0.00022247367,0.00037403897,0.5756318,0.02550047,0.038178008,0.0054987697,0.35173377],"study_design_scores_gemma":[0.000013410676,0.00004363159,0.00015682385,0.000005588546,0.000012068992,0.000024128698,0.000009523006,0.9747395,0.0023587837,0.021866275,0.00076039537,0.000009844245],"about_ca_topic_score_codex":0.0011109821,"about_ca_topic_score_gemma":0.001615955,"teacher_disagreement_score":0.0029470718,"about_ca_system_score_codex":0.0006614582,"about_ca_system_score_gemma":0.0006238935,"threshold_uncertainty_score":0.012250066},"labels":[],"label_agreement":null},{"id":"W2885961833","doi":"10.1177/0165551518787696","title":"Cross-lingual text alignment for fine-grained plagiarism detection","year":2018,"lang":"en","type":"article","venue":"Journal of Information Science","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Institute for Research in Fundamental Sciences; University of Waterloo; University of Tehran","keywords":"Computer science; Plagiarism detection; Similarity (geometry); Natural language processing; Information retrieval; Filter (signal processing); Granularity; Source text; Artificial intelligence; Weighting; Range (aeronautics); Machine translation; Scheme (mathematics); Programming language","score_opus":0.020455169745077065,"score_gpt":0.3085947324421366,"score_spread":0.28813956269705954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885961833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16461974,0.005005793,0.7956952,0.0005226765,0.00027696815,0.0007830434,0.0032889666,0.02155316,0.008254449],"genre_scores_gemma":[0.51392597,0.001198333,0.4709214,0.0001534469,0.0002370457,0.00044530502,0.007149094,0.0011164794,0.004852957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99591124,0.0011080388,0.00053048506,0.001096457,0.0011289074,0.00022490081],"domain_scores_gemma":[0.9910492,0.0028859163,0.0014732595,0.0020045703,0.0022884128,0.0002987006],"candidate_categories":["research_integrity"],"consensus_categories":[],"category_scores_codex":[0.002825356,0.0013261341,0.0016679157,0.011453994,0.0023163336,0.0024146412,0.001322933,0.0014726935,0.0046691317],"category_scores_gemma":[0.013404694,0.0006134453,0.0010188576,0.00902734,0.0007065483,0.0042227698,0.0027367964,0.0014656654,0.005236231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007254222,0.00030729902,0.015187858,0.0012275326,0.00030250984,0.00067925977,0.002468579,0.004574118,0.08810363,0.004773357,0.015079302,0.86657107],"study_design_scores_gemma":[0.0002644359,0.0011123975,0.08714765,0.00034034788,0.00086166296,0.0068295007,0.0042077503,0.56154466,0.20054412,0.032868624,0.10385864,0.00042018134],"about_ca_topic_score_codex":0.0025315904,"about_ca_topic_score_gemma":0.004642366,"teacher_disagreement_score":0.9985273,"about_ca_system_score_codex":0.00075950625,"about_ca_system_score_gemma":0.0016387376,"threshold_uncertainty_score":0.015619814},"labels":[{"model":"gemma","categories":["research_integrity"],"domain":null,"study_design":"simulation_or_modeling","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W2886752202","doi":"10.1139/geomat-2018-0007","title":"A cyclic self-learning Chinese word segmentation for the geoscience domain","year":2018,"lang":"en","type":"article","venue":"GEOMATICA","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Computer science; Context (archaeology); Domain (mathematical analysis); Benchmark (surveying); Natural language processing; Text segmentation; Segmentation; Word (group theory); Artificial intelligence; Information retrieval; Geography; Archaeology; Linguistics; Cartography","score_opus":0.011033181828058703,"score_gpt":0.26550549436588633,"score_spread":0.25447231253782765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886752202","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.094947755,0.0018930422,0.8854444,0.00044589516,0.00033769518,0.0002369759,0.0020090023,0.010524933,0.004160366],"genre_scores_gemma":[0.36677083,0.0009005271,0.6066256,0.00034359138,0.0003354036,0.0003293261,0.012350415,0.0011537084,0.011190692],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99940133,0.00009290831,0.000048348997,0.00028134332,0.00009562544,0.00008032022],"domain_scores_gemma":[0.999041,0.0003933807,0.000044735807,0.0001463988,0.0002850767,0.00008925008],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006259039,0.001242863,0.0011917112,0.0026096958,0.0010348611,0.0008351715,0.0015362544,0.0011310639,0.0050800443],"category_scores_gemma":[0.0015999113,0.00047301422,0.0012236341,0.0027957123,0.00050412974,0.002041991,0.0014545913,0.0012185142,0.002518465],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078568177,0.0003332273,0.002585096,0.00035998554,0.00013869602,0.00029761708,0.00034108467,0.031002225,0.03022595,0.0061318967,0.019970266,0.9078282],"study_design_scores_gemma":[0.00007626951,0.00016179042,0.001461969,0.00002236871,0.00009029553,0.00014743989,0.00016635742,0.97156376,0.011712423,0.0074031106,0.007159408,0.00003485302],"about_ca_topic_score_codex":0.017219344,"about_ca_topic_score_gemma":0.024160618,"teacher_disagreement_score":0.017219344,"about_ca_system_score_codex":0.00073212653,"about_ca_system_score_gemma":0.0022793121,"threshold_uncertainty_score":0.03423822},"labels":[],"label_agreement":null},{"id":"W2886779119","doi":"10.1007/s10844-018-0521-8","title":"Topic and sentiment aware microblog summarization for twitter","year":2018,"lang":"en","type":"article","venue":"Journal of Intelligent Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Microblogging; Social media; Information retrieval; Categorization; Context (archaeology); Multi-document summarization; Sentiment analysis; Process (computing); Task (project management); Representation (politics); World Wide Web; Data science; Natural language processing; Artificial intelligence","score_opus":0.028440984691201352,"score_gpt":0.2675110641839439,"score_spread":0.23907007949274253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886779119","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32804105,0.0071816663,0.60669595,0.0020361205,0.0015982897,0.00074431644,0.020724695,0.025342576,0.007635492],"genre_scores_gemma":[0.6524033,0.0018441293,0.28853577,0.00018880737,0.0017154521,0.0003899742,0.041582108,0.0007864573,0.012553982],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995993,0.00007593518,0.000044835484,0.00009622681,0.00010761034,0.000076136625],"domain_scores_gemma":[0.9990206,0.00028081055,0.00010373842,0.00009949605,0.00041943844,0.00007587958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008231418,0.001083692,0.001035188,0.0033790623,0.0006260596,0.0011067985,0.000594018,0.0006651182,0.0026333542],"category_scores_gemma":[0.0021686605,0.00030033418,0.0008685229,0.0019077283,0.00011392134,0.0012933218,0.00072818465,0.0007238587,0.0028519675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015701,0.00054638606,0.008823134,0.00061207515,0.00038902115,0.00040854208,0.00044397663,0.018421184,0.077559024,0.001260095,0.053128116,0.8368383],"study_design_scores_gemma":[0.00008465598,0.00050023803,0.019140258,0.000049288734,0.00041878526,0.000295105,0.00067904737,0.9171374,0.037067764,0.0032706978,0.021277526,0.000079230565],"about_ca_topic_score_codex":0.003937739,"about_ca_topic_score_gemma":0.008925427,"teacher_disagreement_score":0.003937739,"about_ca_system_score_codex":0.00035193464,"about_ca_system_score_gemma":0.0006037042,"threshold_uncertainty_score":0.008809447},"labels":[],"label_agreement":null},{"id":"W2887005207","doi":"10.18653/v1/w18-3012","title":"Unsupervised Random Walk Sentence Embeddings: A Strong but Simple Baseline","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":95,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Random walk; Sentence; Hyperparameter; Word (group theory); Computer science; Simple (philosophy); Artificial intelligence; Similarity (geometry); Baseline (sea); Natural language processing; Mathematics; Statistics","score_opus":0.024198134676378734,"score_gpt":0.26576507944722183,"score_spread":0.2415669447708431,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2887005207","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025493558,0.0016423838,0.95279974,0.0007460498,0.0004568438,0.00025320062,0.0015579924,0.009972588,0.007077628],"genre_scores_gemma":[0.4326979,0.00090508396,0.5437454,0.0008255049,0.0005957016,0.0005801117,0.006537931,0.0014145216,0.012697743],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976928,0.00083838415,0.00012906363,0.00081647764,0.00039565942,0.0001276644],"domain_scores_gemma":[0.9962949,0.0011751228,0.00022799673,0.0014569041,0.00064528536,0.00019978077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024879612,0.001828741,0.0013496721,0.0015409235,0.00051615643,0.0015681104,0.002243807,0.0023753275,0.006118981],"category_scores_gemma":[0.009252513,0.00064040604,0.0008939114,0.0012031805,0.00060953817,0.0051682508,0.0021363755,0.0023999333,0.006248846],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008098442,0.0010049045,0.0035338718,0.0007051242,0.00044847385,0.00016372888,0.00020878647,0.049566776,0.030929139,0.030288173,0.027168887,0.85517234],"study_design_scores_gemma":[0.000103363054,0.00055419066,0.002333437,0.000066329805,0.00011306185,0.00040300316,0.00005206291,0.91827875,0.014613538,0.042954896,0.020437371,0.000090035],"about_ca_topic_score_codex":0.0014183796,"about_ca_topic_score_gemma":0.0034467115,"teacher_disagreement_score":0.006118981,"about_ca_system_score_codex":0.00055663293,"about_ca_system_score_gemma":0.0011101098,"threshold_uncertainty_score":0.020470023},"labels":[],"label_agreement":null},{"id":"W2887055502","doi":"","title":"Overview of the TREC 2016 Real-Time Summarization Track.","year":2016,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Track (disk drive); Information retrieval; Multi-document summarization; Natural language processing; Operating system","score_opus":0.0542293694693188,"score_gpt":0.2754098846886667,"score_spread":0.22118051521934787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2887055502","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022111543,0.11276645,0.23669279,0.01815191,0.017561253,0.007845086,0.34723586,0.1305811,0.10705406],"genre_scores_gemma":[0.026057877,0.022158554,0.18193594,0.0024723755,0.003704,0.0032939583,0.66582483,0.0059162835,0.08863618],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962171,0.00087844214,0.00044557973,0.00063834625,0.0014156309,0.0004049813],"domain_scores_gemma":[0.9875771,0.0014040495,0.00068935886,0.0012224755,0.00797109,0.0011360004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009083415,0.003022748,0.0023126658,0.010674464,0.002199662,0.004712859,0.0036104373,0.001891015,0.029882034],"category_scores_gemma":[0.009235329,0.0011102845,0.0016611749,0.009430257,0.0004582205,0.0054616393,0.0019398162,0.002588303,0.038880385],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002795532,0.00023671256,0.0006854751,0.0018037214,0.00014742625,0.00008256924,0.00012024494,0.0017480401,0.0134508135,0.00083270745,0.79114455,0.18946818],"study_design_scores_gemma":[0.00023846346,0.0008618749,0.008956989,0.00071170356,0.00041108282,0.0003490705,0.00026991524,0.01500544,0.028222684,0.0032633268,0.94142467,0.00028475147],"about_ca_topic_score_codex":0.041821178,"about_ca_topic_score_gemma":0.07434688,"teacher_disagreement_score":0.041821178,"about_ca_system_score_codex":0.0028850988,"about_ca_system_score_gemma":0.0074805953,"threshold_uncertainty_score":0.09996539},"labels":[],"label_agreement":null},{"id":"W2887344265","doi":"10.1109/icit.2017.12","title":"An Approach to Generic Bengali Text Summarization Using Latent Semantic Analysis","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Future Earth","funders":"","keywords":"Automatic summarization; Bengali; Latent semantic analysis; Computer science; Probabilistic latent semantic analysis; Natural language processing; Multi-document summarization; Artificial intelligence; Semantic analysis (machine learning); Information retrieval; Semantics (computer science)","score_opus":0.08210903640903092,"score_gpt":0.29547942843278235,"score_spread":0.21337039202375144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2887344265","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065590395,0.0011025107,0.9856893,0.00033228498,0.00013500864,0.00017217253,0.00052480696,0.002761855,0.0027231548],"genre_scores_gemma":[0.15019141,0.0017458387,0.8306396,0.00028630372,0.0005341474,0.0004615709,0.0035938523,0.00058747025,0.011959731],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989059,0.00032276913,0.00009296557,0.00028786116,0.00031543625,0.0000750751],"domain_scores_gemma":[0.9987074,0.00030985972,0.00017023791,0.00023616628,0.0005343123,0.000041994666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010559608,0.0011744497,0.00086739444,0.0028675513,0.0008602319,0.0014828242,0.00096021325,0.00067767163,0.002879142],"category_scores_gemma":[0.0024047822,0.0002652844,0.001481113,0.002620794,0.0005001046,0.0016426367,0.0009158489,0.0011694205,0.0019022026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024654006,0.00015466532,0.0011182632,0.00093498133,0.0003277364,0.00027327155,0.0010501652,0.020890104,0.1062572,0.022446236,0.014503657,0.8317972],"study_design_scores_gemma":[0.00012415281,0.000665095,0.008402912,0.00018302735,0.0011699422,0.0010220058,0.0014720743,0.58981395,0.17415299,0.054991093,0.16767737,0.00032536525],"about_ca_topic_score_codex":0.0026331113,"about_ca_topic_score_gemma":0.004633838,"teacher_disagreement_score":0.002879142,"about_ca_system_score_codex":0.00081949116,"about_ca_system_score_gemma":0.0008363174,"threshold_uncertainty_score":0.009631693},"labels":[],"label_agreement":null},{"id":"W2889260178","doi":"10.18653/v1/d18-1544","title":"Grammar Induction with Neural Language Models: An Unusual Replication","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research","funders":"Tencent; Samsung; Nvidia","keywords":"Computer science; Artificial intelligence; Parsing; Grammar; Machine learning; Tree (set theory); Natural language processing; Parse tree; Artificial neural network; Grammar induction; Language model; Tree structure; Rule-based machine translation; Data structure; Linguistics; Programming language","score_opus":0.04716170552813,"score_gpt":0.2846948535745075,"score_spread":0.2375331480463775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2889260178","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14887509,0.0033397947,0.77753425,0.013200065,0.0010669847,0.00030820654,0.0038136453,0.0094268,0.042435177],"genre_scores_gemma":[0.7629815,0.0011448624,0.21495107,0.0019172305,0.00045269236,0.00036526588,0.005969192,0.0021399483,0.010078329],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9885317,0.004527444,0.00056113774,0.0037621951,0.002318229,0.00029931025],"domain_scores_gemma":[0.9600786,0.014078499,0.0006565934,0.021170475,0.0035290138,0.00048687286],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011528927,0.00077750935,0.0010789719,0.0011581525,0.0012124614,0.0034322094,0.0040442515,0.0016224873,0.0053935978],"category_scores_gemma":[0.059868053,0.0006165713,0.0011163923,0.0018210454,0.0026983442,0.009291713,0.0042823176,0.0044952845,0.0048580836],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007907328,0.00053435785,0.024874698,0.0008983818,0.0005223953,0.00092975685,0.0025927278,0.057339497,0.014766355,0.20777449,0.05582144,0.63315517],"study_design_scores_gemma":[0.00014794711,0.00019722567,0.00538893,0.00020241855,0.00012568328,0.0011464648,0.0006920642,0.50946057,0.014069133,0.39297652,0.07545567,0.00013739489],"about_ca_topic_score_codex":0.0051734215,"about_ca_topic_score_gemma":0.004322159,"teacher_disagreement_score":0.9884711,"about_ca_system_score_codex":0.0015069233,"about_ca_system_score_gemma":0.0016033655,"threshold_uncertainty_score":0.0609715},"labels":[],"label_agreement":null},{"id":"W2889675133","doi":"10.18653/v1/d18-1121","title":"Put It Back: Entity Typing with Language Model Enhancement","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Tsinghua University; National Natural Science Foundation of China; National Key Research and Development Program of China; China Association for Science and Technology","keywords":"Computer science; Natural language processing; Language model; Artificial intelligence; Benchmark (surveying); Entity linking; Context (archaeology); Code (set theory); Source code; Typing; Baseline (sea); Information retrieval; Programming language; Speech recognition","score_opus":0.025201139879685613,"score_gpt":0.26459133902248805,"score_spread":0.23939019914280243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2889675133","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02050929,0.0022477473,0.9269727,0.0051206076,0.0017066601,0.0002060151,0.00227325,0.03452135,0.0064423457],"genre_scores_gemma":[0.23651105,0.0014834942,0.7150279,0.0046073655,0.0013439805,0.00030343127,0.012661366,0.004779338,0.023282068],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99679667,0.0013158598,0.00018667952,0.0008701932,0.0006043487,0.0002262153],"domain_scores_gemma":[0.9925196,0.001989892,0.00021503349,0.0036293478,0.0013423545,0.00030384507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004196178,0.002387793,0.0015572928,0.002134424,0.00092742126,0.0027449187,0.0024743502,0.0019163761,0.007294303],"category_scores_gemma":[0.011775368,0.0008744364,0.0024735546,0.0020272618,0.0009018415,0.008927803,0.004360121,0.0060914964,0.01277927],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005025246,0.00042648264,0.0058598462,0.0005074566,0.0003800486,0.00048444068,0.00091507676,0.018953737,0.03022447,0.017930755,0.12479534,0.79901975],"study_design_scores_gemma":[0.00013256232,0.00025246706,0.0027629503,0.00015891965,0.0003609541,0.0014475706,0.00047442518,0.6871107,0.066556096,0.0756208,0.16485946,0.00026308285],"about_ca_topic_score_codex":0.0040593436,"about_ca_topic_score_gemma":0.009562709,"teacher_disagreement_score":0.007294303,"about_ca_system_score_codex":0.00077218417,"about_ca_system_score_gemma":0.0015276864,"threshold_uncertainty_score":0.024401844},"labels":[],"label_agreement":null},{"id":"W2889787757","doi":"10.18653/v1/d18-1259","title":"HotpotQA: A Dataset for Diverse, Explainable Multi-hop Question Answering","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1582,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Office of Naval Research; Defense Advanced Research Projects Agency; Université de Montréal; Nvidia; National Science Foundation","keywords":"Zhàng; Question answering; Computer science; Artificial intelligence; Natural language processing; Information retrieval; History; China; Archaeology","score_opus":0.06108907148323083,"score_gpt":0.31219599013272875,"score_spread":0.2511069186494979,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2889787757","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030940413,0.0045685866,0.012048152,0.0021658759,0.0004489257,0.00085351005,0.9264646,0.015795674,0.0067142346],"genre_scores_gemma":[0.023393963,0.00030716296,0.013626297,0.0004085373,0.000070703616,0.0005091894,0.9589603,0.00023263438,0.002491262],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9976252,0.0006557292,0.00030076777,0.00065859896,0.0005466761,0.00021309919],"domain_scores_gemma":[0.99564016,0.001926495,0.00024218949,0.00092674856,0.00083784905,0.00042644888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019100168,0.0030350233,0.0018565103,0.0043652286,0.0017798989,0.0019126267,0.0043977946,0.005018116,0.013661616],"category_scores_gemma":[0.010889537,0.0006876175,0.0020277055,0.0032733292,0.00071044115,0.003933081,0.0039965827,0.0023434355,0.011121132],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067429716,0.0004935014,0.0057852166,0.0024218368,0.00029271745,0.00053696375,0.00045898315,0.003229815,0.003928535,0.0025752564,0.93685365,0.04274923],"study_design_scores_gemma":[0.0018214989,0.00067668565,0.035531014,0.0008116472,0.0004916476,0.0014935435,0.0025088412,0.09806535,0.010144179,0.02200793,0.82608294,0.00036478907],"about_ca_topic_score_codex":0.03047901,"about_ca_topic_score_gemma":0.06105387,"teacher_disagreement_score":0.03047901,"about_ca_system_score_codex":0.0018902024,"about_ca_system_score_gemma":0.0023384877,"threshold_uncertainty_score":0.0606032},"labels":[],"label_agreement":null},{"id":"W2890236535","doi":"10.18653/v1/d18-1522","title":"Learning Concept Abstractness Using Weak Supervision","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Ilya; Computer science; Empirical research; Natural (archaeology); Artificial intelligence; Natural language processing; Epistemology; Art history; Art; Philosophy; History","score_opus":0.03969259449535188,"score_gpt":0.28793860387932946,"score_spread":0.24824600938397756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890236535","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07472274,0.0031365005,0.9133072,0.0014795377,0.00023021447,0.00014298254,0.0009663174,0.0035266709,0.0024878187],"genre_scores_gemma":[0.7514049,0.0013116265,0.23372324,0.00060499576,0.0006105176,0.00025265643,0.007134802,0.0004822702,0.0044749584],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966498,0.0010292055,0.00022440746,0.0012857978,0.00063945056,0.00017132825],"domain_scores_gemma":[0.9841914,0.009864136,0.0010611333,0.0028715965,0.001432198,0.0005795948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004930339,0.001744603,0.0024585286,0.0027330115,0.0008204292,0.0026054806,0.0025545743,0.0018303023,0.0032785858],"category_scores_gemma":[0.027144114,0.0012360723,0.0017340305,0.001945268,0.0015902709,0.008298863,0.0038387328,0.0053444724,0.0014362993],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017113094,0.00055812835,0.0099821,0.0008106122,0.0007793326,0.00027663342,0.0007814767,0.14869407,0.008814416,0.05656411,0.037080526,0.7339473],"study_design_scores_gemma":[0.00009806247,0.000104679755,0.00095763814,0.00008103092,0.000086449,0.00005773207,0.00005507903,0.84632754,0.0016689446,0.1480842,0.0024542825,0.000024349785],"about_ca_topic_score_codex":0.0025693378,"about_ca_topic_score_gemma":0.0035647762,"teacher_disagreement_score":0.004930339,"about_ca_system_score_codex":0.0010927541,"about_ca_system_score_gemma":0.0013047433,"threshold_uncertainty_score":0.02607441},"labels":[],"label_agreement":null},{"id":"W2890419630","doi":"10.18653/v1/d18-1409","title":"BanditSum: Extractive Summarization as a Contextual Bandit","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":176,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Automatic summarization; Computer science; Artificial intelligence; Reinforcement learning; Context (archaeology); Sequence (biology); Natural language processing; Machine learning","score_opus":0.0206302257562439,"score_gpt":0.26867492822762823,"score_spread":0.24804470247138433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890419630","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011033118,0.00089044374,0.98079085,0.00027496566,0.00011087443,0.00015411737,0.00024504165,0.0048041083,0.0016964953],"genre_scores_gemma":[0.24211027,0.000648341,0.7439243,0.0005434121,0.00033422917,0.0005922306,0.0019391411,0.0008611662,0.009046959],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990802,0.0002912713,0.000067485315,0.00029302042,0.00019683759,0.00007112816],"domain_scores_gemma":[0.99826187,0.00090757204,0.00018304196,0.00025886027,0.00031691368,0.000071748924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017010908,0.0018925347,0.0014843442,0.0013900013,0.0008468312,0.0014373426,0.00228835,0.0015729526,0.0050769625],"category_scores_gemma":[0.005307795,0.00069564267,0.0010197194,0.001232054,0.0007329289,0.0030307064,0.0014595239,0.0021122894,0.0022451538],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040004862,0.00029778067,0.0009509472,0.0004325318,0.00024397482,0.00019465195,0.0003384581,0.35517636,0.018183168,0.011680278,0.009402794,0.602699],"study_design_scores_gemma":[0.00002381968,0.0001295897,0.00017463128,0.000024637917,0.000050047398,0.000042501255,0.00003784212,0.9814341,0.006571828,0.007731069,0.0037625688,0.000017195509],"about_ca_topic_score_codex":0.0033280554,"about_ca_topic_score_gemma":0.008259887,"teacher_disagreement_score":0.0050769625,"about_ca_system_score_codex":0.0010126381,"about_ca_system_score_gemma":0.0012407522,"threshold_uncertainty_score":0.016984165},"labels":[],"label_agreement":null},{"id":"W2890776849","doi":"10.18653/v1/d18-1517","title":"Similar but not the Same: Word Sense Disambiguation Improves Event Detection via Neural Representation Matching","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Université de Montréal","funders":"","keywords":"Computer science; Word (group theory); Artificial intelligence; Natural language processing; Sentence; Event (particle physics); Matching (statistics); Word-sense disambiguation; SemEval; Representation (politics); Artificial neural network; Task (project management); WordNet; Mathematics","score_opus":0.02712744165160249,"score_gpt":0.28412623129732867,"score_spread":0.2569987896457262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890776849","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24532057,0.002467792,0.7362666,0.0012914753,0.00062882557,0.00016201455,0.0008945293,0.008006871,0.004961242],"genre_scores_gemma":[0.8100661,0.0007825896,0.18065186,0.0006780927,0.00030003028,0.000077473414,0.0031139224,0.00026606038,0.0040639495],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99884415,0.00023719521,0.00010258469,0.0005905481,0.00014722829,0.000078244724],"domain_scores_gemma":[0.99824166,0.00075140956,0.00018191268,0.00050749787,0.00023631992,0.00008109897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016152182,0.0010098596,0.0008672374,0.002105086,0.00050233764,0.0012063788,0.00155161,0.0011188698,0.0020315158],"category_scores_gemma":[0.0054850564,0.00027744498,0.0011302426,0.0018611443,0.00061472115,0.0044554593,0.0015368668,0.0013759774,0.0011886837],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005602468,0.0007035318,0.007228755,0.00021822537,0.0002823642,0.00020466659,0.00038384213,0.042048942,0.03318542,0.008661866,0.013417379,0.8931048],"study_design_scores_gemma":[0.00006361033,0.00016806884,0.0056697708,0.000026122363,0.00016406096,0.00023759485,0.0001382325,0.9370316,0.017209526,0.034133438,0.005104663,0.00005330934],"about_ca_topic_score_codex":0.0034471846,"about_ca_topic_score_gemma":0.003487439,"teacher_disagreement_score":0.0034471846,"about_ca_system_score_codex":0.00050579465,"about_ca_system_score_gemma":0.00072579103,"threshold_uncertainty_score":0.00854218},"labels":[],"label_agreement":null},{"id":"W2891212627","doi":"10.1162/coli_a_00333","title":"Introduction to the Special Issue on Language in Social Media: Exploiting Discourse and Other Contextual Information","year":2018,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; University of Ottawa","funders":"","keywords":"Computer science; Social media; Context (archaeology); Perspective (graphical); Process (computing); Interpretation (philosophy); Task (project management); Linguistics; Artificial intelligence; Natural language processing; World Wide Web","score_opus":0.021970997074845143,"score_gpt":0.2991216231871204,"score_spread":0.27715062611227526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891212627","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012125962,0.09839412,0.021117289,0.07413128,0.77328604,0.00020125876,0.0011982165,0.0007549201,0.029704282],"genre_scores_gemma":[0.0038640501,0.061414946,0.0047678486,0.01216064,0.87733656,0.00022008707,0.001598169,0.00074886176,0.03788888],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99783534,0.0004554582,0.00025365205,0.0006019369,0.0006813648,0.0001723608],"domain_scores_gemma":[0.98810136,0.0066246307,0.00064908445,0.0008626998,0.0021605866,0.0016016125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024757995,0.0021023143,0.002479559,0.006466801,0.0024411702,0.008639533,0.002198193,0.0047352365,0.043394927],"category_scores_gemma":[0.009385281,0.00080297224,0.0020097874,0.0047911047,0.002303127,0.00780822,0.004597565,0.007658649,0.020224677],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031800555,0.000065779575,0.00034066636,0.0007645957,0.00003756266,0.00014820574,0.0002521125,0.0001865742,0.00056200684,0.0073251016,0.922327,0.06795875],"study_design_scores_gemma":[0.000005871868,0.000033772256,0.00061526324,0.0005293423,0.000021135711,0.00030222733,0.00013161896,0.00031680128,0.00012936577,0.006135643,0.9917572,0.000021721491],"about_ca_topic_score_codex":0.0009946963,"about_ca_topic_score_gemma":0.0015587424,"teacher_disagreement_score":0.043394927,"about_ca_system_score_codex":0.001756341,"about_ca_system_score_gemma":0.0021562018,"threshold_uncertainty_score":0.14517051},"labels":[],"label_agreement":null},{"id":"W2891369851","doi":"10.18653/v1/d18-1094","title":"A Hierarchical Neural Attention-based Text Classifier","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Interpretability; Computer science; Artificial intelligence; Classifier (UML); Machine learning; Artificial neural network; Deep neural networks; Multiclass classification; Pattern recognition (psychology); Support vector machine","score_opus":0.03405481644673563,"score_gpt":0.2691189879878113,"score_spread":0.23506417154107567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891369851","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09510184,0.0023967219,0.8798359,0.0015833486,0.00059400493,0.00034756988,0.0018067862,0.006176371,0.012157549],"genre_scores_gemma":[0.7143942,0.00082319276,0.25646827,0.0009255448,0.0005075327,0.0003780316,0.003349633,0.00015442232,0.02299921],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999551,0.00006586006,0.000025907477,0.00016806877,0.00011887087,0.00007029218],"domain_scores_gemma":[0.9993685,0.00023764723,0.00005373161,0.00006112605,0.00023105524,0.000048026093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006562909,0.00056138297,0.0006254975,0.0012997059,0.00046571737,0.00070489955,0.0015118439,0.0010003261,0.0037247136],"category_scores_gemma":[0.0019502224,0.0002073344,0.0006769084,0.001164459,0.0002890735,0.0016713947,0.00072284846,0.0010349428,0.0017243951],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002798098,0.00045994026,0.0037049535,0.00020356462,0.00011211884,0.00016735253,0.00012514407,0.06785853,0.0275442,0.008644898,0.030894391,0.86000514],"study_design_scores_gemma":[0.00001549526,0.000058905338,0.0011034677,0.000016875698,0.00003933312,0.0000434965,0.000013251041,0.9884849,0.0036612398,0.004504531,0.0020463762,0.000012147413],"about_ca_topic_score_codex":0.012508892,"about_ca_topic_score_gemma":0.02116753,"teacher_disagreement_score":0.012508892,"about_ca_system_score_codex":0.0012050661,"about_ca_system_score_gemma":0.0013937617,"threshold_uncertainty_score":0.024872184},"labels":[],"label_agreement":null},{"id":"W2891521313","doi":"10.5555/3327757.3327855","title":"Towards Text Generation with Adversarially Learned Neural Outlines","year":2018,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Canadian Institute for Advanced Research; Université de Montréal","funders":"","keywords":"Computer science; Autoregressive model; Artificial intelligence; Generative grammar; Sentence; Generative model; Multinomial distribution; Natural language processing; Latent variable; Machine learning; Prior probability; Bayesian probability; Statistics; Mathematics","score_opus":0.02322451125486854,"score_gpt":0.2462142602693015,"score_spread":0.22298974901443294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891521313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014089036,0.0001486873,0.98273325,0.0002264721,0.000040226492,0.000048768332,0.00009313402,0.0010053429,0.0016150702],"genre_scores_gemma":[0.54574037,0.00032863012,0.44253242,0.00037266457,0.00009081097,0.000286047,0.00077483454,0.00064799975,0.009226215],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995703,0.00018092938,0.000016944361,0.0000957282,0.00010498878,0.00003113123],"domain_scores_gemma":[0.99841535,0.0011359074,0.00009327292,0.0001929012,0.00011394367,0.000048644113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011266142,0.0006700123,0.00049641693,0.00039867457,0.00028063793,0.00070113264,0.0009425131,0.0008351544,0.0032619315],"category_scores_gemma":[0.0038430586,0.00041589997,0.00071047473,0.00029410853,0.0008556736,0.0012880027,0.0010762167,0.0013259054,0.0011721166],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014178918,0.00005854618,0.00074909674,0.00013901675,0.000043685777,0.00023101624,0.00026350567,0.8358728,0.01415297,0.05822529,0.0044018417,0.08572043],"study_design_scores_gemma":[0.000009254623,0.000018063409,0.000047644695,0.0000069494354,0.000004732541,0.000035105008,0.000010147305,0.97820365,0.0032674922,0.01717169,0.001220362,0.0000049362225],"about_ca_topic_score_codex":0.001003141,"about_ca_topic_score_gemma":0.0016764116,"teacher_disagreement_score":0.0032619315,"about_ca_system_score_codex":0.00060620887,"about_ca_system_score_gemma":0.000489894,"threshold_uncertainty_score":0.01091224},"labels":[],"label_agreement":null},{"id":"W2891881175","doi":"10.18653/v1/d18-1343","title":"Multi-Multi-View Learning: Multilingual and Multi-Representation Entity Typing","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Computer science; Natural language processing; Embedding; Representation (politics); Artificial intelligence; Entity linking; Context (archaeology); German; Information retrieval; Linguistics; Knowledge base","score_opus":0.08456332633316738,"score_gpt":0.3519646892348669,"score_spread":0.2674013629016995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891881175","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11790104,0.0051302337,0.8142987,0.0014242688,0.00046128707,0.0004641882,0.034336865,0.019879514,0.006103993],"genre_scores_gemma":[0.3958084,0.0009322164,0.46976086,0.00078237854,0.00022052569,0.0003618491,0.12791201,0.0008432651,0.0033784946],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99672455,0.0008799363,0.0002859458,0.0013672208,0.00052487757,0.00021746411],"domain_scores_gemma":[0.9935847,0.0025065574,0.00030479222,0.002632855,0.0006536734,0.0003174365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036031753,0.0018491361,0.0017514393,0.0038442893,0.0011240033,0.0027404728,0.0038664548,0.002211984,0.003914638],"category_scores_gemma":[0.011457764,0.00066947617,0.0021287864,0.0050342404,0.00074496755,0.006900279,0.004833982,0.0034166717,0.0023161385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010920656,0.0013734059,0.02656425,0.0013025294,0.0008547859,0.0008160432,0.0007246211,0.09036851,0.008615662,0.008710789,0.12271936,0.736858],"study_design_scores_gemma":[0.0001757823,0.00025316663,0.0060088886,0.00017641309,0.0002683588,0.00095835,0.00058245397,0.90062463,0.013863974,0.035752498,0.041232917,0.00010253898],"about_ca_topic_score_codex":0.008800288,"about_ca_topic_score_gemma":0.015986718,"teacher_disagreement_score":0.008800288,"about_ca_system_score_codex":0.0012109326,"about_ca_system_score_gemma":0.0012877703,"threshold_uncertainty_score":0.019055665},"labels":[],"label_agreement":null},{"id":"W2892043975","doi":"10.18653/v1/d18-1502","title":"Dual Fixed-Size Ordinally Forgetting Encoding (FOFE) for Competitive Neural Language Models","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Forgetting; Dual (grammatical number); Computer science; Encoding (memory); Perplexity; Artificial neural network; Word (group theory); Language model; Artificial intelligence; Deep neural networks; Mathematics; Linguistics","score_opus":0.028342806514413312,"score_gpt":0.2765021994500761,"score_spread":0.24815939293566278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2892043975","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016203517,0.00065149204,0.9790036,0.00027035017,0.00012772522,0.000042356925,0.00014514492,0.002203492,0.0013523757],"genre_scores_gemma":[0.57161444,0.0005710866,0.42008865,0.00057204795,0.00023274303,0.00020130939,0.0008669202,0.0004159993,0.0054367823],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918073,0.0002526617,0.00007514783,0.00020779223,0.00018715617,0.00009652155],"domain_scores_gemma":[0.9977113,0.0011562625,0.00014212664,0.0005005182,0.00038884123,0.000100933896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016120571,0.0013433766,0.0011816877,0.0008051885,0.00041303292,0.0014424043,0.0033267383,0.0015749983,0.004443673],"category_scores_gemma":[0.007813981,0.0004736766,0.00096319633,0.000816435,0.0007568384,0.004377658,0.0013911268,0.0029152231,0.0014612942],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040857823,0.00026186413,0.0018053795,0.0002723099,0.00016833596,0.00032230117,0.00021583415,0.32406983,0.008215021,0.04694773,0.00671095,0.61060184],"study_design_scores_gemma":[0.000012607914,0.000038432263,0.00008179296,0.000010580006,0.000016214037,0.000055673205,0.000009513786,0.98405564,0.0016130799,0.013131924,0.0009605913,0.0000139389695],"about_ca_topic_score_codex":0.00536914,"about_ca_topic_score_gemma":0.0077661527,"teacher_disagreement_score":0.00536914,"about_ca_system_score_codex":0.00091278896,"about_ca_system_score_gemma":0.0010915675,"threshold_uncertainty_score":0.014865577},"labels":[],"label_agreement":null},{"id":"W2892355797","doi":"10.7939/dvn/10968","title":"WNED datasets and results","year":2017,"lang":"en","type":"dataset","venue":"Borealis","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Geography; Archaeology; Cartography","score_opus":0.03943210719498486,"score_gpt":0.30497854617797987,"score_spread":0.265546438982995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2892355797","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011855725,0.00024086033,0.0005854908,0.00024244677,0.00026165147,0.00015290621,0.99226224,0.0024167728,0.0026519813],"genre_scores_gemma":[0.0006394185,0.000074922995,0.0010730603,0.00007483495,0.00001559015,0.00019368362,0.99675757,0.00010056956,0.0010703328],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99655414,0.0006199587,0.00045661107,0.0008660911,0.0010396627,0.00046354692],"domain_scores_gemma":[0.99571675,0.0008007834,0.00023874953,0.0016272984,0.0012015228,0.00041500974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036967972,0.0048332317,0.0023320909,0.006363779,0.0021404366,0.0028172012,0.005577169,0.0029446098,0.046349302],"category_scores_gemma":[0.010023881,0.00086694525,0.002672598,0.007171003,0.0010824822,0.002950452,0.003687364,0.0035420964,0.07921862],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014149823,0.00014498465,0.0007607299,0.0003928041,0.000048401158,0.00003759122,0.000017184635,0.00084424193,0.00021424175,0.0005532404,0.9900787,0.0067664166],"study_design_scores_gemma":[0.000591113,0.00014504895,0.0049925246,0.00035044955,0.00010982397,0.0002656424,0.00032033515,0.0063960045,0.0026184306,0.004467156,0.9796402,0.00010324271],"about_ca_topic_score_codex":0.028727246,"about_ca_topic_score_gemma":0.062745236,"teacher_disagreement_score":0.046349302,"about_ca_system_score_codex":0.0031834412,"about_ca_system_score_gemma":0.0038088502,"threshold_uncertainty_score":0.15505391},"labels":[],"label_agreement":null},{"id":"W2893437765","doi":"10.48550/arxiv.1809.11086","title":"Learning Recurrent Binary/Ternary Weights","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Ternary operation; Recurrent neural network; Speedup; Binary number; Application-specific integrated circuit; Inference; Sequence (biology); Throughput; Identification (biology); Parallel computing; Algorithm; Computer engineering; Computer hardware; Artificial intelligence; Artificial neural network; Arithmetic; Programming language; Operating system","score_opus":0.07755215751959647,"score_gpt":0.19364457622969694,"score_spread":0.11609241871010047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2893437765","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07873135,0.00029554145,0.9112539,0.00025570786,0.00019484754,0.00004692464,0.0003689874,0.0048518004,0.004000978],"genre_scores_gemma":[0.69267344,0.00028903186,0.29845187,0.00024676096,0.000109932174,0.00011003905,0.0008776092,0.0003614661,0.0068798065],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99969184,0.000031767362,0.000025550853,0.00010460335,0.0000955678,0.000050754756],"domain_scores_gemma":[0.9993649,0.00020487454,0.00008149672,0.00012572865,0.00018862735,0.00003445737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00042476697,0.0008424846,0.000504949,0.0005991811,0.00028746185,0.0008981627,0.0011601212,0.00069533713,0.0035973147],"category_scores_gemma":[0.003301126,0.00040695156,0.0004908141,0.0005945703,0.00032801743,0.0017683381,0.00075269584,0.00095114636,0.0013812444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033664447,0.00014364552,0.0029838863,0.00025234916,0.0001384826,0.0003204112,0.00018025262,0.29193458,0.07082332,0.026237499,0.010156912,0.596492],"study_design_scores_gemma":[0.000009156668,0.00002648779,0.0002848893,0.000008513183,0.00002101164,0.00005155676,0.000015917816,0.9724139,0.018037014,0.00753801,0.0015835359,0.000010053152],"about_ca_topic_score_codex":0.003599811,"about_ca_topic_score_gemma":0.0063051055,"teacher_disagreement_score":0.003599811,"about_ca_system_score_codex":0.0005321299,"about_ca_system_score_gemma":0.0008283668,"threshold_uncertainty_score":0.012034237},"labels":[],"label_agreement":null},{"id":"W2894762547","doi":"10.1145/3209280.3229114","title":"Improving Short Text Clustering by Similarity Matrix Sparsification","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Cluster analysis; Similarity (geometry); Computer science; Artificial intelligence; Embedding; Word (group theory); Correlation clustering; Document clustering; Pattern recognition (psychology); Data mining; Mathematics; Image (mathematics)","score_opus":0.028969456681981722,"score_gpt":0.27476297035351493,"score_spread":0.2457935136715332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2894762547","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057479214,0.00095855416,0.9362311,0.00027354062,0.00014648713,0.00012594929,0.00042550956,0.0030105961,0.0013490297],"genre_scores_gemma":[0.35233918,0.00084779935,0.6367681,0.0002501371,0.00026918866,0.00022726573,0.0038160689,0.0006755386,0.0048067835],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987832,0.00029292743,0.00011058251,0.00033979045,0.00038379864,0.000089731024],"domain_scores_gemma":[0.9959126,0.0017809281,0.0003748833,0.0009523958,0.0008071004,0.00017202998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011062151,0.001212809,0.0014319288,0.0016798434,0.0007382945,0.0011972327,0.0011184834,0.001115,0.0028187742],"category_scores_gemma":[0.00844011,0.00034635165,0.00092463533,0.0021684337,0.0007111351,0.0028777393,0.0014993673,0.0012465412,0.0032613147],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081635034,0.00031332625,0.0030970883,0.00072594034,0.0002506904,0.0003131515,0.0004967307,0.21159531,0.11998863,0.008317976,0.012642496,0.6414423],"study_design_scores_gemma":[0.000031836986,0.0001368038,0.0010184159,0.000015252816,0.000040987325,0.00021259717,0.00016894772,0.94948596,0.036923464,0.008405699,0.0035279964,0.000032072065],"about_ca_topic_score_codex":0.002474062,"about_ca_topic_score_gemma":0.0042729424,"teacher_disagreement_score":0.0028187742,"about_ca_system_score_codex":0.00042060227,"about_ca_system_score_gemma":0.0008179648,"threshold_uncertainty_score":0.009429693},"labels":[],"label_agreement":null},{"id":"W2894995323","doi":"10.1007/978-3-030-01716-3_32","title":"Coherence-Based Automated Essay Scoring Using Self-attention","year":2018,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Feature engineering; Coherence (philosophical gambling strategy); Artificial intelligence; Representation (politics); Feature (linguistics); Machine learning; Natural language processing; Deep learning; Linguistics; Statistics","score_opus":0.028839341247385297,"score_gpt":0.26396261349139655,"score_spread":0.23512327224401125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2894995323","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20789474,0.0016605032,0.74527764,0.0004697706,0.00081536913,0.00068067183,0.0022644952,0.026040616,0.014896271],"genre_scores_gemma":[0.72882193,0.0002586738,0.24685168,0.00012908214,0.0005878033,0.00041620605,0.005447138,0.0008931307,0.016594268],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99533147,0.0016571367,0.0003463577,0.0010277461,0.0012850397,0.00035226595],"domain_scores_gemma":[0.983085,0.007687558,0.0010682897,0.0018468216,0.00561134,0.0007010039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00375505,0.0011363112,0.0014263422,0.003487217,0.00073271844,0.0020098018,0.0014080958,0.0009285229,0.007987394],"category_scores_gemma":[0.016135672,0.00038841087,0.0005152752,0.0020213644,0.0002775447,0.002207391,0.0026718762,0.0011232429,0.006341803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005701786,0.0003812031,0.009994492,0.00021895474,0.00015407919,0.00008090641,0.00029289885,0.0045147217,0.020300766,0.000994231,0.020515248,0.9419823],"study_design_scores_gemma":[0.00018799542,0.00077230873,0.034948085,0.00006243277,0.0002425137,0.0003221968,0.00059217133,0.9125438,0.029886154,0.008741661,0.011577121,0.00012363977],"about_ca_topic_score_codex":0.001561701,"about_ca_topic_score_gemma":0.0035917931,"teacher_disagreement_score":0.007987394,"about_ca_system_score_codex":0.0003650643,"about_ca_system_score_gemma":0.0007593195,"threshold_uncertainty_score":0.026720464},"labels":[],"label_agreement":null},{"id":"W2895258120","doi":"10.5121/ijma.2016.8401","title":"A Document Exploring System on LDA Topic Model for Wikipedia Articles","year":2016,"lang":"en","type":"article","venue":"The International journal of Multimedia & Its Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Acadia University","funders":"","keywords":"Information retrieval; Topic model; Computer science; World Wide Web","score_opus":0.08308563270336597,"score_gpt":0.30061917942811583,"score_spread":0.21753354672474987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2895258120","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03278338,0.0010213108,0.78683317,0.00044339226,0.00018933519,0.00080765726,0.0041586543,0.16753198,0.0062311585],"genre_scores_gemma":[0.1490008,0.00070400175,0.82362974,0.00024161667,0.00015589816,0.0009980928,0.009835999,0.0036476643,0.011786119],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99940705,0.00012099399,0.00005923323,0.000204882,0.00016770515,0.00004006453],"domain_scores_gemma":[0.99890435,0.00037748486,0.000070535636,0.00018474899,0.00036942592,0.00009356348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001301206,0.0009251601,0.00084641366,0.0036380084,0.0010438716,0.0016693877,0.0010870586,0.0008100704,0.004790857],"category_scores_gemma":[0.0033087286,0.00045865838,0.00111117,0.0022319963,0.00029429872,0.0028704267,0.0013172859,0.0007452985,0.00475215],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008664949,0.0004288097,0.0060687265,0.0010402446,0.00033563736,0.00054759096,0.0018863379,0.0072352844,0.05649553,0.005774066,0.07370633,0.8456151],"study_design_scores_gemma":[0.00026787174,0.0005074178,0.007147138,0.00015373665,0.00039951503,0.0018472481,0.0011606339,0.69946945,0.09177859,0.012312293,0.18455148,0.00040455876],"about_ca_topic_score_codex":0.0042130845,"about_ca_topic_score_gemma":0.004627297,"teacher_disagreement_score":0.004790857,"about_ca_system_score_codex":0.00046334034,"about_ca_system_score_gemma":0.0011083157,"threshold_uncertainty_score":0.016026974},"labels":[],"label_agreement":null},{"id":"W2895279374","doi":"10.1162/coli_a_00363","title":"Scalable Micro-planned Generation of Discourse from Structured Data","year":2019,"lang":"en","type":"preprint","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Interpretability; Natural language processing; Scalability; Artificial intelligence; Natural language generation; Sentence; Pipeline (software); Paragraph; Fluency; Text simplification; Natural language understanding; Robustness (evolution); Natural language; Programming language; Database","score_opus":0.08137583821686917,"score_gpt":0.32304697739436655,"score_spread":0.24167113917749738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2895279374","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011597212,0.00023911726,0.9495867,0.0004327829,0.0001281203,0.00031311784,0.004045073,0.031025302,0.0026325872],"genre_scores_gemma":[0.098622896,0.00019295917,0.88170624,0.00016219438,0.00007755972,0.00038980832,0.013734699,0.0020190943,0.0030945688],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986174,0.00043286625,0.00010352433,0.00048448678,0.0003039159,0.000057732675],"domain_scores_gemma":[0.99565303,0.0026105454,0.00022161052,0.0007322887,0.00065607607,0.00012636877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018243723,0.0010039036,0.0007394824,0.0016972892,0.00068932574,0.0016535218,0.0016916131,0.000788645,0.0076516536],"category_scores_gemma":[0.009486726,0.0005961796,0.0012221023,0.0013254888,0.00070986757,0.0026309418,0.002192564,0.0012214081,0.0043866206],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006422142,0.00029728469,0.0031811034,0.0014926012,0.00015361143,0.0010173338,0.002867715,0.03825135,0.054153807,0.05084812,0.07385127,0.7732434],"study_design_scores_gemma":[0.0001494971,0.00014909543,0.00092405453,0.0001088072,0.000075998076,0.0003946891,0.00075096876,0.7789135,0.075950645,0.060524378,0.08197053,0.00008784291],"about_ca_topic_score_codex":0.0026794751,"about_ca_topic_score_gemma":0.004420707,"teacher_disagreement_score":0.0076516536,"about_ca_system_score_codex":0.0007892889,"about_ca_system_score_gemma":0.0019335772,"threshold_uncertainty_score":0.025597334},"labels":[],"label_agreement":null},{"id":"W2895332542","doi":"10.1145/3209280.3229101","title":"OurDirection","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Recurrent neural network; Encoder; Domain (mathematical analysis); Artificial intelligence; Artificial neural network; Open domain; Question answering; Operating system","score_opus":0.023879334514023368,"score_gpt":0.25892729411985416,"score_spread":0.2350479596058308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2895332542","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011809234,0.002464754,0.7236203,0.023623586,0.0046587954,0.0009661182,0.0066270395,0.020605018,0.20562522],"genre_scores_gemma":[0.18302983,0.0022918845,0.6211761,0.008805165,0.0021034235,0.001153355,0.022187905,0.0042009614,0.15505138],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99493706,0.002145125,0.00016354094,0.0012655973,0.0010958244,0.0003927913],"domain_scores_gemma":[0.99485946,0.00110584,0.00013567945,0.0019936485,0.001153394,0.00075210276],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.005031149,0.0017320218,0.00071619387,0.0012652134,0.0022929092,0.006387888,0.0035864792,0.0028359324,0.07045866],"category_scores_gemma":[0.011549725,0.00053603953,0.0013685713,0.0008293896,0.0017572833,0.010371583,0.0070366603,0.003822415,0.03394595],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041548364,0.00042744886,0.0022674603,0.00046795706,0.000054999666,0.00025077502,0.0011263851,0.00637164,0.0046195243,0.34959912,0.24398454,0.3904147],"study_design_scores_gemma":[0.00012522035,0.00015663686,0.0006325024,0.00024667243,0.000038511094,0.00040865317,0.0005975995,0.053321782,0.004967106,0.16964966,0.76976424,0.00009146415],"about_ca_topic_score_codex":0.0069914,"about_ca_topic_score_gemma":0.008327973,"teacher_disagreement_score":0.92954135,"about_ca_system_score_codex":0.0018852197,"about_ca_system_score_gemma":0.004330348,"threshold_uncertainty_score":0.23570764},"labels":[],"label_agreement":null},{"id":"W2896065871","doi":"10.1109/ijcnn.2018.8489224","title":"Words Are Not Temporal Sequences of Characters","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Recurrent neural network; Artificial intelligence; Generative grammar; Word (group theory); Natural language processing; Architecture; Language model; Implementation; Component (thermodynamics); Invariant (physics); Natural language; Generative model; Encoding (memory); Position (finance); Artificial neural network; Speech recognition; Linguistics; Programming language; Mathematics","score_opus":0.046210009076187386,"score_gpt":0.26922645979226517,"score_spread":0.22301645071607779,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896065871","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17473777,0.0016395332,0.7884154,0.0010262994,0.0008290347,0.00020717135,0.0037128374,0.0020554315,0.027376495],"genre_scores_gemma":[0.87425226,0.0011427287,0.10102359,0.0002118138,0.00021437867,0.0001859833,0.0031038662,0.00032991252,0.019535461],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9996338,0.00007229402,0.000026630203,0.0001556739,0.000081900165,0.000029770872],"domain_scores_gemma":[0.9989586,0.00036610244,0.00019012499,0.00027795276,0.00017489132,0.000032192187],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00032332088,0.00048377438,0.00033803258,0.00056034507,0.00031603654,0.0014560283,0.0007664447,0.00062609126,0.0054522115],"category_scores_gemma":[0.004523993,0.00037459147,0.00038952846,0.0012555526,0.0006380031,0.0044441056,0.00046753587,0.00094491564,0.002142604],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006832059,0.0001402946,0.010661571,0.0009970238,0.00019840884,0.00077035424,0.0018035388,0.07323283,0.065303646,0.42819613,0.013232683,0.40478033],"study_design_scores_gemma":[0.000066710345,0.00030607905,0.009151669,0.0002257008,0.00021195451,0.0011573008,0.0006391603,0.5175864,0.031189002,0.30144364,0.13792677,0.000095669406],"about_ca_topic_score_codex":0.0019745785,"about_ca_topic_score_gemma":0.0037715556,"teacher_disagreement_score":0.0054522115,"about_ca_system_score_codex":0.00035390843,"about_ca_system_score_gemma":0.00057302864,"threshold_uncertainty_score":0.018239498},"labels":[],"label_agreement":null},{"id":"W2896529804","doi":"10.1007/s12083-018-0687-4","title":"Toward citation recommender systems considering the article impact in the extended nearby citation network","year":2018,"lang":"en","type":"article","venue":"Peer-to-Peer Networking and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Citation; Computer science; Impact factor; Metadata; Bibliometrics; Similarity (geometry); Information retrieval; Recommender system; Citation impact; Data science; Empirical research; Baseline (sea); Data mining; World Wide Web; Artificial intelligence; Statistics; Mathematics","score_opus":0.0670989261935125,"score_gpt":0.3148766972970984,"score_spread":0.2477777711035859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896529804","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21403031,0.00841737,0.766113,0.0026089654,0.00049803447,0.00014938576,0.00077160104,0.00048813314,0.006923194],"genre_scores_gemma":[0.90838975,0.0041546235,0.07750421,0.00027691707,0.0015148869,0.00013846638,0.0008348505,0.00009041261,0.007095969],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99831223,0.0006172648,0.00009911053,0.00049466005,0.00033281863,0.00014389656],"domain_scores_gemma":[0.9889737,0.007439677,0.0009161084,0.00054320914,0.0016863404,0.00044097845],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0032203423,0.0009450714,0.0022248568,0.0035149993,0.0012223427,0.0032874688,0.0029022265,0.0027860592,0.0017484516],"category_scores_gemma":[0.017990846,0.00079669914,0.001171957,0.0054910295,0.0007795328,0.004292889,0.0015343693,0.0017846606,0.0006760074],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041214327,0.0005344776,0.030733135,0.0006438157,0.0009573132,0.0005654729,0.00065603515,0.70431894,0.0038711221,0.06400047,0.010249181,0.18305789],"study_design_scores_gemma":[0.00001032111,0.000028440183,0.0010979647,0.000017593571,0.000084503205,0.00004121414,0.00003427895,0.98540777,0.00018606216,0.012223073,0.00085133204,0.000017363776],"about_ca_topic_score_codex":0.011557649,"about_ca_topic_score_gemma":0.014788665,"teacher_disagreement_score":0.996485,"about_ca_system_score_codex":0.0011862963,"about_ca_system_score_gemma":0.0014474309,"threshold_uncertainty_score":0.02298075},"labels":[],"label_agreement":null},{"id":"W2896581050","doi":"10.1109/ijcnn.2018.8489194","title":"Using Deep Learning to Recommend Discussion Threads to Users in an Online Forum","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Set (abstract data type); Recall; Conversation; Artificial neural network; Test set; Topic model; Artificial intelligence; Probabilistic logic; Sample (material); Social media; Ideal (ethics); Machine learning; F1 score; Precision and recall; Data set; Test (biology); World Wide Web","score_opus":0.09314019694162755,"score_gpt":0.34527129089458264,"score_spread":0.2521310939529551,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2896581050","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5245891,0.0016588389,0.45641837,0.0016269851,0.00033814952,0.00053273654,0.0016212125,0.005788346,0.007426187],"genre_scores_gemma":[0.8946495,0.00028723702,0.09919301,0.00014502188,0.00008093775,0.0001747732,0.0015262319,0.0000598237,0.0038833858],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990006,0.00040001996,0.00007028055,0.0002698935,0.00013904445,0.00012008051],"domain_scores_gemma":[0.9954397,0.0028040388,0.0003657157,0.00032328858,0.00083963905,0.00022757603],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029234109,0.0011659397,0.0004922769,0.002438081,0.0006864053,0.0013132691,0.0010134064,0.0012333955,0.001271367],"category_scores_gemma":[0.0093868105,0.00055750215,0.00066823506,0.0013190409,0.00030791125,0.002489114,0.0010462741,0.0017329769,0.0008323602],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001015783,0.0016875006,0.048731342,0.0004008342,0.00031408292,0.00021795822,0.001957957,0.14495274,0.009798571,0.003679233,0.011423436,0.7758205],"study_design_scores_gemma":[0.00003191961,0.000108103726,0.002603026,0.000028133469,0.000035464695,0.000026952372,0.00013178178,0.98948073,0.0023856047,0.0037496446,0.0014015007,0.000017029022],"about_ca_topic_score_codex":0.013886086,"about_ca_topic_score_gemma":0.028030058,"teacher_disagreement_score":0.013886086,"about_ca_system_score_codex":0.0015898462,"about_ca_system_score_gemma":0.0013912978,"threshold_uncertainty_score":0.02761054},"labels":[],"label_agreement":null},{"id":"W2898396042","doi":"10.1109/asonam.2018.8508612","title":"Implicit Entity Linking Through Ad-Hoc Retrieval","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Entity linking; Information retrieval; Knowledge base; Set (abstract data type); Process (computing); Representation (politics); Space (punctuation); Baseline (sea); Natural language processing; World Wide Web; Programming language","score_opus":0.035632039234080135,"score_gpt":0.28899289894091085,"score_spread":0.2533608597068307,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898396042","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025342967,0.0024124542,0.96029663,0.0005362372,0.0001920755,0.0005290917,0.0009165599,0.0042609773,0.0055130357],"genre_scores_gemma":[0.31260794,0.0028489593,0.6549037,0.00047944373,0.00058422325,0.00064510613,0.005926983,0.0006781821,0.02132546],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951379,0.0018643029,0.0004049307,0.0011459728,0.0011435851,0.00030342618],"domain_scores_gemma":[0.98820156,0.0048611322,0.0008931134,0.004473252,0.0014175493,0.00015344116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045207385,0.0018846589,0.0021746412,0.006572613,0.0017143682,0.0041024047,0.0044470243,0.0028126235,0.006977312],"category_scores_gemma":[0.018228523,0.00091258105,0.0015090341,0.008246671,0.0015724255,0.011112346,0.0043526096,0.0019551031,0.0060804696],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047172533,0.0006115625,0.003808599,0.0008713215,0.00031925004,0.00042863225,0.00083824195,0.06431429,0.018339494,0.03449991,0.023770714,0.85172623],"study_design_scores_gemma":[0.00020532176,0.0004311968,0.002782296,0.0001532159,0.00047734714,0.0014398897,0.0007904736,0.8100335,0.043020714,0.07625436,0.064220235,0.00019139954],"about_ca_topic_score_codex":0.0050847293,"about_ca_topic_score_gemma":0.005991636,"teacher_disagreement_score":0.006977312,"about_ca_system_score_codex":0.0012036086,"about_ca_system_score_gemma":0.0027214692,"threshold_uncertainty_score":0.023908257},"labels":[],"label_agreement":null},{"id":"W2898487021","doi":"10.48550/arxiv.1810.10641","title":"Predicting the Semantic Textual Similarity with Siamese CNN and LSTM","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Natural language processing; Semantic similarity; Similarity (geometry); Artificial intelligence; Context (archaeology); Convolution (computer science); Convolutional neural network; Recurrent neural network; Artificial neural network; Image (mathematics)","score_opus":0.05907066756665994,"score_gpt":0.17951226178288918,"score_spread":0.12044159421622924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898487021","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43343365,0.0020789357,0.54242295,0.0011524814,0.0005671842,0.00023173133,0.002822589,0.00832997,0.008960601],"genre_scores_gemma":[0.9091177,0.00039308742,0.08234031,0.0001907775,0.00021328546,0.00008977551,0.0027812591,0.0001428277,0.004731058],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996532,0.0000630436,0.000025011974,0.00015319574,0.000059421083,0.000046071673],"domain_scores_gemma":[0.99939394,0.00024417232,0.00007308445,0.00006449901,0.00018286477,0.000041488693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00060746947,0.0009733889,0.0005636063,0.0013915979,0.0002951563,0.000819358,0.00080341153,0.001012127,0.0026437645],"category_scores_gemma":[0.0024857398,0.00028237927,0.00072999496,0.001330282,0.00027827334,0.0020210268,0.00053365825,0.0008412424,0.0013207411],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086589373,0.0005910674,0.011567893,0.00035255507,0.00043661246,0.0005343387,0.00023878673,0.15742183,0.05497075,0.0064989924,0.019192407,0.747329],"study_design_scores_gemma":[0.000008611745,0.00003863085,0.0011117607,0.000005717151,0.00002575933,0.000039189646,0.00001833974,0.992088,0.0034687072,0.0026936978,0.00049413496,0.000007403505],"about_ca_topic_score_codex":0.007460216,"about_ca_topic_score_gemma":0.011958263,"teacher_disagreement_score":0.007460216,"about_ca_system_score_codex":0.0008720962,"about_ca_system_score_gemma":0.0005871744,"threshold_uncertainty_score":0.0148335695},"labels":[],"label_agreement":null},{"id":"W2898607795","doi":"10.2196/medinform.9965","title":"Clinical Named Entity Recognition From Chinese Electronic Health Records via Machine Learning Methods","year":2018,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Conditional random field; Artificial intelligence; Natural language processing; F1 score; Machine learning; Named-entity recognition; Test set; Benchmark (surveying); Precision and recall; Health records; Deep learning; Information retrieval; Unstructured data; Task (project management); Data mining; Big data; Health care","score_opus":0.04074273956746675,"score_gpt":0.4129983582792558,"score_spread":0.372255618711789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898607795","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19368218,0.0067403237,0.74531436,0.002982043,0.00076901447,0.0013702956,0.027478369,0.014325292,0.00733811],"genre_scores_gemma":[0.48504883,0.0020869565,0.45261434,0.0006222734,0.00040835832,0.00077839097,0.05387522,0.00018524719,0.0043803896],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99737394,0.0005485427,0.00042663002,0.0011390572,0.00034658457,0.00016528579],"domain_scores_gemma":[0.9954118,0.002327,0.0006195296,0.0007545461,0.0007970575,0.00009008124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024784955,0.0013062528,0.0008932409,0.0051464476,0.0008324764,0.0013274874,0.0018078362,0.0012704416,0.0021014942],"category_scores_gemma":[0.0090753455,0.00032315784,0.0014841309,0.004643659,0.0005607238,0.0029305622,0.0013236974,0.0011696585,0.0015853319],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004030159,0.00026041045,0.017069923,0.00087039935,0.00023264633,0.0015516421,0.0006136375,0.048462596,0.008216026,0.006262977,0.026521688,0.8895351],"study_design_scores_gemma":[0.00008243589,0.00012474744,0.013784575,0.00017383574,0.00027629727,0.0008878089,0.000418288,0.92966336,0.017154248,0.013532564,0.02377579,0.00012602667],"about_ca_topic_score_codex":0.01654226,"about_ca_topic_score_gemma":0.012885066,"teacher_disagreement_score":0.01654226,"about_ca_system_score_codex":0.0014303188,"about_ca_system_score_gemma":0.002965925,"threshold_uncertainty_score":0.03289193},"labels":[],"label_agreement":null},{"id":"W2898749396","doi":"10.18653/v1/w18-5620","title":"Listwise temporal ordering of events in clinical notes","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institutes of Health; University of Toronto; Mayo Clinic","keywords":"Timeline; Computer science; Baseline (sea); Ranking (information retrieval); Downstream (manufacturing); Natural language processing; Information retrieval; Statistics; Engineering; Mathematics; Geology","score_opus":0.08409123596687995,"score_gpt":0.3602126343501249,"score_spread":0.27612139838324495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898749396","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26441345,0.008069482,0.6493684,0.0021043154,0.00076970353,0.0008112155,0.055572044,0.009907012,0.008984503],"genre_scores_gemma":[0.6438411,0.0017603665,0.28580213,0.00029867422,0.0005492061,0.00048363197,0.0632127,0.0006052449,0.0034470165],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941062,0.0015716641,0.0010363105,0.0008998373,0.0020245914,0.00036142033],"domain_scores_gemma":[0.9684507,0.018433161,0.0042264676,0.001901506,0.0058946796,0.001093491],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044883722,0.0009102369,0.0007168951,0.0065922947,0.00067807606,0.0021012232,0.0010772978,0.00087832403,0.0039101257],"category_scores_gemma":[0.03893181,0.00026549798,0.0007233646,0.0051582917,0.00032667583,0.0031821053,0.001209645,0.0012025669,0.0018532826],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002879021,0.00060677424,0.16709457,0.0033387006,0.0005703787,0.0008393985,0.001034928,0.07833157,0.0308038,0.018092912,0.056070562,0.6403374],"study_design_scores_gemma":[0.00023642568,0.0016988617,0.11839334,0.0005074583,0.0007529668,0.0032859708,0.00127521,0.6951242,0.05211311,0.071400985,0.054873545,0.0003380111],"about_ca_topic_score_codex":0.006341571,"about_ca_topic_score_gemma":0.012218919,"teacher_disagreement_score":0.0065922947,"about_ca_system_score_codex":0.0009264148,"about_ca_system_score_gemma":0.002523711,"threshold_uncertainty_score":0.023737073},"labels":[],"label_agreement":null},{"id":"W2899231639","doi":"10.18653/v1/w19-4103","title":"Augmenting Neural Response Generation with Context-Aware Topical Attention","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Encoder; Context (archaeology); Natural language generation; Attention network; Artificial intelligence; Similarity (geometry); Artificial neural network; Quality (philosophy); Recurrent neural network; Machine learning; Natural language","score_opus":0.05465600436238829,"score_gpt":0.27336988893461783,"score_spread":0.21871388457222954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899231639","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11703648,0.00069459603,0.8600474,0.0007169306,0.0002726513,0.00025043235,0.0010116828,0.014748621,0.0052212533],"genre_scores_gemma":[0.7409104,0.00026696196,0.24441987,0.0005110024,0.00015228524,0.00035193772,0.0027577754,0.00070953305,0.009920207],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989029,0.0005299979,0.00003781628,0.00029167125,0.00015204221,0.00008545421],"domain_scores_gemma":[0.99672,0.0020515407,0.00014429264,0.0003817239,0.0005533107,0.00014913929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019178378,0.0010839398,0.00063178455,0.00058745616,0.00029021257,0.00059732795,0.0014771818,0.0010126738,0.0028718044],"category_scores_gemma":[0.0078046224,0.00037210147,0.0004991274,0.00041914952,0.00039773562,0.0012669156,0.0014242119,0.001528606,0.0023647274],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009187655,0.0006801501,0.0056051128,0.00045070998,0.00016876336,0.0003938628,0.0008351433,0.16913398,0.097114,0.0057991454,0.016111987,0.70278835],"study_design_scores_gemma":[0.000030175726,0.00010736429,0.00048380974,0.000012050869,0.000024860656,0.00006736143,0.00006722948,0.9752864,0.01877411,0.0028218448,0.0023061116,0.000018597288],"about_ca_topic_score_codex":0.0035942376,"about_ca_topic_score_gemma":0.006658213,"teacher_disagreement_score":0.0035942376,"about_ca_system_score_codex":0.0005812811,"about_ca_system_score_gemma":0.0009013763,"threshold_uncertainty_score":0.010142624},"labels":[],"label_agreement":null},{"id":"W2899336413","doi":"10.18653/v1/d18-1524","title":"InferLite: Simple Universal Sentence Representations from Natural Language Inference Data","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada)","funders":"","keywords":"Computer science; Inference; Natural language processing; Sentence; Artificial intelligence; Word (group theory); Simple (philosophy); Natural language; Context (archaeology); Task (project management); Word order; Language model; Natural language understanding; Linguistics","score_opus":0.051124446523659044,"score_gpt":0.3373849179116872,"score_spread":0.2862604713880281,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899336413","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024058266,0.00054445904,0.9248676,0.0007363465,0.00024714737,0.00020041806,0.009200154,0.037335776,0.002809877],"genre_scores_gemma":[0.3078435,0.00067636656,0.6405308,0.0007400185,0.00027207346,0.0005563361,0.03956433,0.002082883,0.007733625],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99939704,0.00012831592,0.000041531493,0.00025866277,0.0001173426,0.000057089856],"domain_scores_gemma":[0.9983224,0.00066844653,0.00011793956,0.00064941903,0.000166899,0.000074865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013531039,0.0012194522,0.00069786416,0.0012328323,0.0004213995,0.0014900756,0.0024908325,0.0012249057,0.0086864745],"category_scores_gemma":[0.0073383194,0.00068589457,0.0013485253,0.0009613645,0.00064108515,0.006464776,0.0022506819,0.0026893164,0.0041728034],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005939342,0.00045926808,0.0044286004,0.0008194759,0.0003303009,0.00037605638,0.000598262,0.086055405,0.018824864,0.051570717,0.08616691,0.7497761],"study_design_scores_gemma":[0.000036310114,0.000068042005,0.00070361246,0.00004929503,0.000045095847,0.00015215017,0.000065248154,0.9195204,0.008589823,0.058638975,0.012091742,0.000039310948],"about_ca_topic_score_codex":0.0036332775,"about_ca_topic_score_gemma":0.008159586,"teacher_disagreement_score":0.0086864745,"about_ca_system_score_codex":0.00082505785,"about_ca_system_score_gemma":0.0012456285,"threshold_uncertainty_score":0.029059172},"labels":[],"label_agreement":null},{"id":"W2899463504","doi":"10.18653/v1/w18-5618","title":"In-domain Context-aware Token Embeddings Improve Biomedical Named Entity Recognition","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"Natural Sciences and Engineering Research Council of Canada; Ministère de la Défense Nationale; Genome British Columbia; Genome Canada","keywords":"Computer science; Named-entity recognition; Pipeline (software); Security token; Entity linking; Natural language processing; Domain (mathematical analysis); Artificial intelligence; Context (archaeology); Biomedical text mining; Named entity; Task (project management); Natural language; Information retrieval; Text mining; Knowledge base; Programming language","score_opus":0.01883654162191673,"score_gpt":0.26487343160900306,"score_spread":0.24603688998708634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899463504","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2775572,0.0051136315,0.6673477,0.002373665,0.0010776774,0.00019299002,0.0084743295,0.030109027,0.0077537545],"genre_scores_gemma":[0.7191991,0.0012009821,0.25140613,0.0005669306,0.00030667082,0.00012732694,0.019292088,0.00069893774,0.0072017633],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904925,0.00032158973,0.00009801219,0.00032390837,0.000105492836,0.00010176892],"domain_scores_gemma":[0.99691534,0.001542443,0.000267277,0.00074174756,0.00040959753,0.00012356813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014829667,0.00113325,0.0007163317,0.0018114331,0.00044613905,0.0013192194,0.0010583068,0.001305509,0.0033727428],"category_scores_gemma":[0.0062876134,0.00032845116,0.0009027975,0.0019862792,0.00034408623,0.005336463,0.0015838424,0.0019606433,0.003422463],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007072108,0.00076545455,0.013177976,0.0004849543,0.00030923297,0.00034410882,0.0004431379,0.073302574,0.018911686,0.0078570545,0.03769725,0.84599936],"study_design_scores_gemma":[0.000053881377,0.00019570146,0.0032622386,0.000060435967,0.00015176817,0.00024906942,0.00023066218,0.94742703,0.018200006,0.016011564,0.014097717,0.00005993346],"about_ca_topic_score_codex":0.0034742202,"about_ca_topic_score_gemma":0.006855193,"teacher_disagreement_score":0.0034742202,"about_ca_system_score_codex":0.00060948153,"about_ca_system_score_gemma":0.000999495,"threshold_uncertainty_score":0.011282921},"labels":[],"label_agreement":null},{"id":"W2899501643","doi":"10.18653/v1/p19-1386","title":"The KnowRef Coreference Corpus: Removing Gender and Number Cues for Difficult Pronominal Anaphora Resolution","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research; Microsoft Research","keywords":"Coreference; Computer science; Antecedent (behavioral psychology); Anaphora (linguistics); Benchmark (surveying); Natural language processing; Artificial intelligence; Task (project management); Context (archaeology); Feature (linguistics); Resolution (logic); Linguistics; Psychology","score_opus":0.03347918295935006,"score_gpt":0.25855318998511284,"score_spread":0.22507400702576277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899501643","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38711163,0.017294014,0.3371564,0.0049546286,0.002434832,0.0017863773,0.15444438,0.038516484,0.056301303],"genre_scores_gemma":[0.34598735,0.0019155674,0.31345966,0.0011823052,0.00040811725,0.002268999,0.32070747,0.0041618138,0.009908681],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99497354,0.0023130823,0.00042697656,0.0012325877,0.0008328252,0.0002209178],"domain_scores_gemma":[0.9862185,0.0074444166,0.00045429077,0.0037885304,0.0018168434,0.0002773509],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041818093,0.0015835365,0.0013376882,0.004343906,0.003953681,0.001799829,0.0028206285,0.0031761276,0.0077477535],"category_scores_gemma":[0.020109367,0.00064790994,0.0007672533,0.0042288518,0.0016055508,0.004415453,0.0038780442,0.0030645106,0.0043158517],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000995483,0.0010815866,0.0059044654,0.006975337,0.00032345363,0.0017003675,0.005258283,0.018145386,0.04109532,0.018279659,0.43371806,0.4665226],"study_design_scores_gemma":[0.00066629803,0.00055936736,0.024345262,0.0008994968,0.0002447513,0.0042247213,0.0056083393,0.14088638,0.093329534,0.032412395,0.6964645,0.00035903894],"about_ca_topic_score_codex":0.0063892286,"about_ca_topic_score_gemma":0.0139208855,"teacher_disagreement_score":0.0077477535,"about_ca_system_score_codex":0.00093246537,"about_ca_system_score_gemma":0.0020303319,"threshold_uncertainty_score":0.025918782},"labels":[],"label_agreement":null},{"id":"W2900341241","doi":"10.18653/v1/w18-6250","title":"UBC-NLP at IEST 2018: Learning Implicit Emotion With an Ensemble of Language Models","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Task (project management); Artificial intelligence; Ranking (information retrieval); Natural language processing; Baseline (sea); Language model; Training set; Ensemble forecasting; Machine learning; Speech recognition","score_opus":0.0231619943106445,"score_gpt":0.25262752461390825,"score_spread":0.22946553030326375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2900341241","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15016073,0.0048969965,0.5514673,0.0058333655,0.0048246332,0.0015095855,0.041692257,0.20778045,0.0318348],"genre_scores_gemma":[0.39501354,0.0013350899,0.42632526,0.0019815406,0.0009624961,0.0018439032,0.12668231,0.008202339,0.03765348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99865144,0.00034284257,0.0000672403,0.0004918523,0.0002850313,0.00016160501],"domain_scores_gemma":[0.9984975,0.0005007674,0.00004906849,0.0004311888,0.00040135885,0.00012020721],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002582859,0.0025230988,0.0014831419,0.00094347587,0.0010872777,0.002297549,0.0018207151,0.00185219,0.013302529],"category_scores_gemma":[0.0065874313,0.00050047995,0.0010447955,0.00097538426,0.0005534695,0.003955734,0.0027646192,0.004760976,0.013613874],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010475218,0.000997604,0.0049464465,0.0008532368,0.00044225363,0.000887851,0.0008288757,0.034894917,0.05041641,0.004315273,0.32231948,0.5780502],"study_design_scores_gemma":[0.00026422608,0.00048621467,0.0041601416,0.00012746736,0.00017108122,0.00059721357,0.0005088241,0.82564056,0.04909197,0.009972407,0.10880991,0.00016994245],"about_ca_topic_score_codex":0.007262626,"about_ca_topic_score_gemma":0.009576358,"teacher_disagreement_score":0.013302529,"about_ca_system_score_codex":0.0010104058,"about_ca_system_score_gemma":0.0016385363,"threshold_uncertainty_score":0.044501424},"labels":[],"label_agreement":null},{"id":"W2902138027","doi":"10.1007/978-3-030-17705-8_9","title":"Generating Responses Expressing Emotion in an Open-Domain Dialogue System","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Sentence; Encoder; Artificial intelligence; Natural language processing; Artificial neural network; Open domain; Domain (mathematical analysis); Speech recognition; Machine learning; Question answering","score_opus":0.04062409163093213,"score_gpt":0.27503194008160775,"score_spread":0.23440784845067564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902138027","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3378932,0.00015525198,0.6366428,0.00044320195,0.00023727867,0.0005660147,0.00046762818,0.015820434,0.0077742357],"genre_scores_gemma":[0.8085793,0.000080345664,0.17986177,0.0002639269,0.00008391143,0.00041408782,0.00077913736,0.0007916628,0.009145785],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99890614,0.00064400485,0.000044158085,0.00020557373,0.00013898623,0.00006117404],"domain_scores_gemma":[0.9959598,0.0033377495,0.00008759723,0.00016503395,0.00025861379,0.00019119444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001522805,0.0008422132,0.0007447004,0.0002666614,0.00040682335,0.0011977662,0.0011624099,0.001417676,0.007654037],"category_scores_gemma":[0.0052391947,0.00026877405,0.00039828557,0.0002029767,0.00049010676,0.00093516486,0.0016030448,0.0007148485,0.0022914372],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0066304137,0.0020000467,0.004146796,0.0011761654,0.00024426277,0.0026932887,0.014978886,0.03719766,0.4970275,0.009288179,0.01704847,0.4075683],"study_design_scores_gemma":[0.0008124764,0.0019788058,0.0043734238,0.000113941875,0.00025839105,0.0009471231,0.0037497994,0.80005914,0.15502688,0.013522671,0.018947642,0.00020977293],"about_ca_topic_score_codex":0.00030451198,"about_ca_topic_score_gemma":0.00029556433,"teacher_disagreement_score":0.007654037,"about_ca_system_score_codex":0.00022188808,"about_ca_system_score_gemma":0.00020097024,"threshold_uncertainty_score":0.02560532},"labels":[],"label_agreement":null},{"id":"W2902360757","doi":"10.18653/v1/w18-6518","title":"Automatic Opinion Question Generation","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada; University of Lethbridge","keywords":"Computer science; Question answering; Natural language processing; Artificial intelligence; Sequence (biology); Mechanism (biology); Information retrieval","score_opus":0.044301483495684336,"score_gpt":0.29861374351005493,"score_spread":0.2543122600143706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902360757","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043662433,0.0012275161,0.9397904,0.0013017363,0.0003670473,0.00046108913,0.0016481482,0.0066588954,0.004882752],"genre_scores_gemma":[0.45721954,0.00049564784,0.52523655,0.00066863047,0.0006575986,0.00052176695,0.01032073,0.0007119949,0.0041675745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99486333,0.0028276301,0.00022538222,0.00089651626,0.00091423723,0.00027288837],"domain_scores_gemma":[0.9872951,0.0077220183,0.0005345768,0.0012142065,0.0029115772,0.000322576],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047172965,0.0010717589,0.0011134985,0.0020248205,0.00074293074,0.0013106201,0.0019057029,0.0017306661,0.0050879046],"category_scores_gemma":[0.020688096,0.00032064965,0.0011930554,0.0011635736,0.00048123093,0.0028888776,0.0017569521,0.0013795091,0.003172283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045169357,0.00057369884,0.00583486,0.0010722439,0.00014886583,0.00063507824,0.0013981084,0.021569077,0.069596395,0.024365041,0.05874861,0.8156064],"study_design_scores_gemma":[0.000098954246,0.00022618863,0.0022184905,0.00006431748,0.000104343024,0.0006476931,0.0003205789,0.88403666,0.04160557,0.04417136,0.026455877,0.000049940583],"about_ca_topic_score_codex":0.0010573167,"about_ca_topic_score_gemma":0.0010663159,"teacher_disagreement_score":0.0050879046,"about_ca_system_score_codex":0.00060562155,"about_ca_system_score_gemma":0.0007524251,"threshold_uncertainty_score":0.024947703},"labels":[],"label_agreement":null},{"id":"W2902761513","doi":"10.12688/gatesopenres.12891.2","title":"Automated verbal autopsy classification: using one-against-all ensemble method and Naïve Bayes classifier","year":2019,"lang":"en","type":"preprint","venue":"Gates Open Research","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto; Centre for Global Health Research; St. Michael's Hospital; Toronto Metropolitan University","funders":"Bill and Melinda Gates Foundation","keywords":"Verbal autopsy; Artificial intelligence; Naive Bayes classifier; Classifier (UML); Pattern recognition (psychology); Bayes' theorem; Computer science; Natural language processing; Bayesian probability; Support vector machine; Medicine; Pathology; Cause of death","score_opus":0.40473298567594607,"score_gpt":0.4866043839533529,"score_spread":0.08187139827740686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902761513","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36345336,0.0048740054,0.6143865,0.0010522808,0.0009893345,0.0007306641,0.002261095,0.006190636,0.006062172],"genre_scores_gemma":[0.77160966,0.00066006,0.21800606,0.0004419247,0.0004126819,0.0002959961,0.0050922614,0.00016133967,0.0033200283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967276,0.0010747546,0.00038456355,0.0008355124,0.0006516304,0.00032593473],"domain_scores_gemma":[0.9953282,0.0022102997,0.0002735478,0.00043260484,0.0015876357,0.00016771007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055914866,0.0022496574,0.0026623455,0.0034192754,0.0012697235,0.0014237135,0.0018782421,0.0016864071,0.0013786335],"category_scores_gemma":[0.008550521,0.00039591113,0.0016634642,0.0015399584,0.00037224626,0.0010888309,0.0010374893,0.0015967807,0.0008606386],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084651896,0.00066575524,0.03742037,0.00024170763,0.0007179717,0.0003335432,0.00020656752,0.1264913,0.0043364344,0.0010506402,0.009734801,0.8179545],"study_design_scores_gemma":[0.000042183507,0.0002623943,0.0048888945,0.00005947546,0.00023853523,0.00016872422,0.00011746059,0.98620427,0.003715985,0.0023560058,0.0019021976,0.00004383037],"about_ca_topic_score_codex":0.011960197,"about_ca_topic_score_gemma":0.009832357,"teacher_disagreement_score":0.011960197,"about_ca_system_score_codex":0.00068190147,"about_ca_system_score_gemma":0.0017049912,"threshold_uncertainty_score":0.029570937},"labels":[],"label_agreement":null},{"id":"W2903988268","doi":"10.1609/aaai.v33i01.33017031","title":"Fast PMI-Based Word Embedding with Efficient Use of Unobserved Patterns","year":2019,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Pointwise mutual information; Word (group theory); Computer science; Word embedding; Artificial intelligence; Embedding; Similarity (geometry); Pointwise; Margin (machine learning); Natural language processing; Semantic similarity; Kernel (algebra); Latent semantic analysis; Vocabulary; Mutual information; Machine learning; Mathematics; Linguistics","score_opus":0.09438551158186712,"score_gpt":0.27984865010700816,"score_spread":0.18546313852514104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2903988268","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012228127,0.00015983662,0.9857136,0.000091130554,0.000040627576,0.00005337879,0.00015187419,0.001100438,0.00046108259],"genre_scores_gemma":[0.23204987,0.00024946246,0.7599489,0.00012223935,0.00010496128,0.00038946516,0.0022249704,0.00040578732,0.0045044054],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991047,0.00024248079,0.000073011186,0.00025812583,0.0002449748,0.0000766409],"domain_scores_gemma":[0.99792624,0.000972032,0.0001857069,0.00043788395,0.00039621364,0.000081826416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007979315,0.0010504489,0.0011609121,0.0014349355,0.00047851048,0.0008888202,0.0017386219,0.00091517554,0.0029966417],"category_scores_gemma":[0.0046129427,0.000559232,0.0009612756,0.0019974473,0.0006901576,0.003090106,0.0018785961,0.001802975,0.0023193439],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024330743,0.00018338382,0.0020059994,0.00022580034,0.000101283156,0.00012928457,0.00023845017,0.114112884,0.01634219,0.021735027,0.006426656,0.83825564],"study_design_scores_gemma":[0.000013045441,0.000046332218,0.0003334196,0.000006829109,0.000008240259,0.00006667343,0.000029575065,0.9820326,0.0033144134,0.012932009,0.0012042364,0.0000126483565],"about_ca_topic_score_codex":0.0022147188,"about_ca_topic_score_gemma":0.003985354,"teacher_disagreement_score":0.0029966417,"about_ca_system_score_codex":0.00057187903,"about_ca_system_score_gemma":0.0010532576,"threshold_uncertainty_score":0.010024786},"labels":[],"label_agreement":null},{"id":"W2904366287","doi":"10.1145/3293339.3293345","title":"Question-Question Similarity in Online Forums","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Artificial intelligence; Similarity (geometry); Representation (politics); Context (archaeology); Natural language processing; Sentence; Recall; Task (project management); Deep learning; Recurrent neural network; Meaning (existential); Convolutional neural network; Semantic similarity; Question answering; Term (time); Feature learning; Machine learning; Artificial neural network","score_opus":0.027474850297985114,"score_gpt":0.30600385017260134,"score_spread":0.2785289998746162,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2904366287","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49993628,0.00891205,0.4642186,0.0013368113,0.00078767724,0.0009809118,0.0075500715,0.0076635177,0.008614028],"genre_scores_gemma":[0.8854571,0.000491803,0.10037232,0.00016418865,0.00028035883,0.00022669596,0.009445194,0.0001683009,0.003394033],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99455076,0.001883791,0.0005637087,0.0015566068,0.0010940093,0.00035108314],"domain_scores_gemma":[0.9867018,0.0075189006,0.0015243854,0.0014434821,0.00223905,0.00057228905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005215119,0.00091422687,0.0014139721,0.007396909,0.0009653042,0.0019855837,0.0012985602,0.0018403209,0.0034035782],"category_scores_gemma":[0.027014276,0.0003093309,0.0012472097,0.0030638669,0.000538885,0.006855543,0.002437384,0.00095622917,0.0012331865],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00176835,0.0012907074,0.08242739,0.0029588433,0.00048811163,0.0012722743,0.0043159565,0.03765762,0.032724068,0.021265663,0.019318113,0.7945129],"study_design_scores_gemma":[0.00013368476,0.00087724795,0.06760316,0.0003058178,0.00033707093,0.0016129805,0.0030056976,0.76237935,0.03494577,0.08448696,0.044126786,0.00018543251],"about_ca_topic_score_codex":0.0029979884,"about_ca_topic_score_gemma":0.0029800565,"teacher_disagreement_score":0.007396909,"about_ca_system_score_codex":0.0013808997,"about_ca_system_score_gemma":0.0010786072,"threshold_uncertainty_score":0.02758056},"labels":[],"label_agreement":null},{"id":"W2904760604","doi":"10.1609/aaai.v33i01.33016188","title":"Enriching Word Embeddings with a Regressor Instead of Labeled Corpora","year":2019,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; Vetenskapsrådet","keywords":"Computer science; Natural language processing; Artificial intelligence; Word (group theory); Lexicon; Task (project management); Embedding; Sentence; SemEval; Word embedding; Analogy; Binary classification; Linguistics; Support vector machine","score_opus":0.05268464246031737,"score_gpt":0.27048664703602743,"score_spread":0.21780200457571006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2904760604","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02897238,0.00038552127,0.9553441,0.00042474284,0.00017120737,0.00013627586,0.0017453479,0.009922393,0.0028981154],"genre_scores_gemma":[0.20797457,0.0005746546,0.7695757,0.00038282276,0.00016244533,0.00055043556,0.011510326,0.0015272405,0.007741793],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981438,0.00073891843,0.00013863113,0.0005662267,0.0003244473,0.00008800097],"domain_scores_gemma":[0.99451435,0.001860121,0.00049700576,0.0018737887,0.0011190397,0.00013574677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023195194,0.0021184848,0.0010455163,0.0023644362,0.00060435414,0.001643724,0.0013741087,0.0010292256,0.0039025669],"category_scores_gemma":[0.013444678,0.0007583957,0.0007617848,0.0024628278,0.0007702289,0.0077308803,0.0028812536,0.002251119,0.006540991],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034512184,0.00070792023,0.010921096,0.0006991458,0.00022359299,0.00025131518,0.0006889406,0.034517735,0.062244665,0.016469534,0.024316017,0.84861493],"study_design_scores_gemma":[0.00007628378,0.0004660993,0.0043343864,0.00019953032,0.00018457515,0.0004953316,0.00062427035,0.7724243,0.08025922,0.06029651,0.08046087,0.00017851495],"about_ca_topic_score_codex":0.0018293122,"about_ca_topic_score_gemma":0.0060157212,"teacher_disagreement_score":0.0039025669,"about_ca_system_score_codex":0.0004905641,"about_ca_system_score_gemma":0.0012940346,"threshold_uncertainty_score":0.013055325},"labels":[],"label_agreement":null},{"id":"W2904940351","doi":"10.1609/aaai.v33i01.33013526","title":"Improved Knowledge Graph Embedding Using Background Taxonomic Information","year":2019,"lang":"en","type":"preprint","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Embedding; Knowledge graph; Computer science; Theoretical computer science; Graph; Mathematics; Artificial intelligence","score_opus":0.13480551769492677,"score_gpt":0.32662901701316743,"score_spread":0.19182349931824066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2904940351","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031466894,0.00031507429,0.96222365,0.00051418936,0.00004168981,0.00007230395,0.00089088693,0.0023884734,0.0020868797],"genre_scores_gemma":[0.3969902,0.00061885593,0.5858138,0.0002761034,0.000103682156,0.00016095924,0.0071004573,0.00071956095,0.008216351],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987079,0.00035216755,0.00007785201,0.0004485621,0.00031878435,0.000094692004],"domain_scores_gemma":[0.99486876,0.0020363126,0.00032903932,0.0021364796,0.00046063584,0.00016886381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012298257,0.0010075892,0.0010210735,0.0017358069,0.0005942046,0.0019793767,0.001653702,0.0013577932,0.0039986027],"category_scores_gemma":[0.008706959,0.00058675423,0.0015086908,0.0022183687,0.00084136904,0.008258321,0.0030657477,0.002488,0.001408617],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030959735,0.00050412543,0.0026045016,0.00058662007,0.00019733375,0.00045620298,0.00089537253,0.28558362,0.015476768,0.12686518,0.020782946,0.5457378],"study_design_scores_gemma":[0.000016026623,0.00003078648,0.0002583795,0.00002670331,0.000037777845,0.00009915344,0.00008832752,0.9094244,0.0024222168,0.083538845,0.004041754,0.00001559341],"about_ca_topic_score_codex":0.0042843614,"about_ca_topic_score_gemma":0.006619258,"teacher_disagreement_score":0.0042843614,"about_ca_system_score_codex":0.00087700685,"about_ca_system_score_gemma":0.0011232108,"threshold_uncertainty_score":0.013376653},"labels":[],"label_agreement":null},{"id":"W2906152891","doi":"10.1162/tacl_a_00254","title":"Analysis Methods in Neural Language Processing: A Survey","year":2019,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":486,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Categorization; Artificial neural network; Field (mathematics); Feature (linguistics); Artificial intelligence; Point (geometry); Data science; Natural language processing; Machine learning; Linguistics","score_opus":0.029637724967615017,"score_gpt":0.3494461539490562,"score_spread":0.3198084289814412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2906152891","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004142226,0.507405,0.47500902,0.0034250747,0.0008518541,0.00016084961,0.00053162826,0.00094141375,0.0075329836],"genre_scores_gemma":[0.09496226,0.54976726,0.3416665,0.0014634586,0.0041957134,0.0005342282,0.0016020169,0.00064952386,0.005159032],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962387,0.0011962429,0.00042164308,0.0006630212,0.0013786331,0.0001018874],"domain_scores_gemma":[0.989958,0.0071156304,0.00037459744,0.0005820197,0.001865727,0.00010402895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052472623,0.0011753333,0.0015038105,0.006815145,0.0005551925,0.0033278705,0.0021171414,0.0013811911,0.0033386108],"category_scores_gemma":[0.0135822585,0.0006397034,0.0013988793,0.008449441,0.0011510465,0.004993781,0.0012643527,0.0019748725,0.00211179],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000088158355,0.00011778804,0.0029042237,0.0036488199,0.00025758648,0.00007869989,0.00022937522,0.0055933446,0.0012275785,0.049447425,0.016287267,0.9201197],"study_design_scores_gemma":[0.00007161051,0.00021386819,0.008603497,0.00473148,0.00041577066,0.0011577583,0.00066121324,0.24980146,0.0070504462,0.31669512,0.41039133,0.00020645266],"about_ca_topic_score_codex":0.0026796907,"about_ca_topic_score_gemma":0.0015001078,"teacher_disagreement_score":0.006815145,"about_ca_system_score_codex":0.0013949091,"about_ca_system_score_gemma":0.0017379277,"threshold_uncertainty_score":0.027750552},"labels":[],"label_agreement":null},{"id":"W2906208572","doi":"10.1609/aaai.v33i01.330110075","title":"Sequence to Sequence Learning for Query Expansion","year":2019,"lang":"en","type":"preprint","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Sequence (biology); Query expansion; Information retrieval; Set (abstract data type); Sentence; Space (punctuation); Natural language processing; Artificial intelligence","score_opus":0.20417178667479077,"score_gpt":0.351188536485249,"score_spread":0.14701674981045823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2906208572","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054371975,0.001389447,0.93303937,0.0005260146,0.00013828986,0.00032553176,0.000533298,0.0047913822,0.004884673],"genre_scores_gemma":[0.5372406,0.0007714456,0.45114315,0.00053283473,0.00021679325,0.0004310758,0.0022625413,0.00040739827,0.006994136],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987997,0.0004297777,0.000091129594,0.00035291846,0.00023580842,0.000090810165],"domain_scores_gemma":[0.9974935,0.0014486546,0.00013677483,0.00035431693,0.0004907264,0.000076031465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018062192,0.0009301286,0.00067731366,0.0013589568,0.00032425867,0.0005957025,0.0011688849,0.0010023768,0.006547138],"category_scores_gemma":[0.007335757,0.00029104244,0.0006503379,0.0012523028,0.0005614608,0.0035563107,0.000991131,0.0013582879,0.002148574],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051530695,0.0005087734,0.0021425325,0.0005658045,0.00010282361,0.00016848498,0.00035166965,0.11609196,0.039124057,0.019414632,0.0120046465,0.80900925],"study_design_scores_gemma":[0.00003219419,0.00019499903,0.00043101632,0.0000191403,0.000023158727,0.000112152724,0.00006015969,0.96751904,0.009801306,0.017214283,0.0045781424,0.000014338349],"about_ca_topic_score_codex":0.0034809539,"about_ca_topic_score_gemma":0.004570568,"teacher_disagreement_score":0.006547138,"about_ca_system_score_codex":0.00084123516,"about_ca_system_score_gemma":0.0008958209,"threshold_uncertainty_score":0.021902323},"labels":[],"label_agreement":null},{"id":"W2907750764","doi":"10.4000/books.aaccademia.4862","title":"The Perfect Recipe: Add SUGAR, Add Data","year":2018,"lang":"en","type":"book-chapter","venue":"Accademia University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Research Council Canada; Università degli Studi di Napoli Federico II","keywords":"Task (project management); Recipe; Computer science; Training set; Set (abstract data type); Utterance; Representation (politics); Artificial intelligence; Natural language processing; Robot; Data set; Speech recognition; Engineering; Programming language; Geography","score_opus":0.07694988937194383,"score_gpt":0.24531033904953495,"score_spread":0.16836044967759112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2907750764","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012244821,0.0034719007,0.4663812,0.008759539,0.0073512625,0.0005656086,0.056176074,0.11869336,0.32635623],"genre_scores_gemma":[0.0727874,0.0025262388,0.39668193,0.0023562163,0.0011051449,0.00057680305,0.07907213,0.04944439,0.39544976],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995679,0.00009655638,0.000018078812,0.00010933135,0.0001649499,0.000043111006],"domain_scores_gemma":[0.99930155,0.00023964953,0.000018307774,0.00023018928,0.00013335886,0.000076986784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008548904,0.0011475054,0.000619541,0.00079388754,0.00081601617,0.0021612737,0.0012634383,0.0009953054,0.17686749],"category_scores_gemma":[0.0035841337,0.00051217905,0.00061233394,0.0010780884,0.0005003201,0.0041032443,0.0021226339,0.0015222112,0.1333664],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019300333,0.00007619018,0.00036190767,0.0002921181,0.00002562778,0.00015954416,0.00033773258,0.0017373896,0.003583556,0.013411992,0.7608245,0.21899644],"study_design_scores_gemma":[0.000028791887,0.00003945681,0.00030899126,0.000045418714,0.0000100977,0.0001386737,0.0001548515,0.0058360146,0.0033254528,0.012592431,0.97748715,0.000032608277],"about_ca_topic_score_codex":0.0032592013,"about_ca_topic_score_gemma":0.0076160114,"teacher_disagreement_score":0.17686749,"about_ca_system_score_codex":0.00045662528,"about_ca_system_score_gemma":0.0005830687,"threshold_uncertainty_score":0.59168065},"labels":[],"label_agreement":null},{"id":"W2907786475","doi":"10.1109/icosc.2019.8665673","title":"Improving Tree-LSTM with Tree Attention","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Tree (set theory); Computer science; Tree structure; Artificial intelligence; Dependency (UML); Encoding (memory); Task (project management); Sentence; Decision tree model; Natural language processing; Machine learning; Theoretical computer science; Decision tree; Algorithm; Mathematics; Binary tree","score_opus":0.02104007912626795,"score_gpt":0.22960562832488846,"score_spread":0.2085655491986205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2907786475","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048521098,0.0024000565,0.9370226,0.00062909164,0.0003587202,0.000064806925,0.00034911075,0.006456523,0.004198006],"genre_scores_gemma":[0.7158479,0.0014626634,0.26889178,0.0010938129,0.00035008654,0.00015589395,0.0016989358,0.00080744206,0.0096914675],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99956733,0.00010931981,0.000025602063,0.00013508018,0.00009869848,0.00006396027],"domain_scores_gemma":[0.99897265,0.0006240929,0.00005535765,0.00009751846,0.00020898753,0.000041330644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00099053,0.0011356609,0.00092952093,0.0006865361,0.0003328412,0.00073719665,0.0013084051,0.0015127115,0.0028824168],"category_scores_gemma":[0.003820612,0.00039723355,0.0007381999,0.0011078423,0.0003041333,0.003539248,0.0012065937,0.0017531018,0.0012232809],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000275296,0.00023556188,0.0011072085,0.0003216898,0.00019721693,0.00018876932,0.00027295132,0.32449654,0.020244371,0.013092774,0.017607579,0.62196004],"study_design_scores_gemma":[0.000010543804,0.000026012876,0.00011540893,0.00000671997,0.000023263196,0.000019542385,0.000007820828,0.9919694,0.0019027264,0.0050824415,0.0008312701,0.0000049230553],"about_ca_topic_score_codex":0.007962869,"about_ca_topic_score_gemma":0.011931965,"teacher_disagreement_score":0.007962869,"about_ca_system_score_codex":0.00081846496,"about_ca_system_score_gemma":0.00096834806,"threshold_uncertainty_score":0.01583308},"labels":[],"label_agreement":null},{"id":"W2908434667","doi":"10.4000/books.aaccademia.4613","title":"A Kernel-based Approach for Irony and Sarcasm Detection in Italian","year":2018,"lang":"en","type":"book-chapter","venue":"Accademia University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Research Council Canada; Università degli Studi di Napoli Federico II","keywords":"Sarcasm; Irony; Support vector machine; Artificial intelligence; Task (project management); Computer science; Kernel (algebra); Context (archaeology); Natural language processing; Machine learning; Linguistics; Mathematics; Engineering; History; Philosophy","score_opus":0.03612086206838847,"score_gpt":0.21594969685849486,"score_spread":0.1798288347901064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2908434667","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35703504,0.004719195,0.58336115,0.0015002025,0.000687824,0.00039513916,0.0025660053,0.022702737,0.027032698],"genre_scores_gemma":[0.8037152,0.0008187314,0.16747968,0.00016535257,0.00027033038,0.00016177546,0.004285266,0.0005233154,0.022580294],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991881,0.00021782372,0.000048236812,0.00022917426,0.00019032526,0.00012632136],"domain_scores_gemma":[0.9992347,0.0002869107,0.00008182578,0.00012967188,0.00021362322,0.00005338518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009976584,0.0010473876,0.0005760336,0.0015576954,0.00058731093,0.001217048,0.0006851432,0.0008834271,0.0035938553],"category_scores_gemma":[0.002839955,0.00031641306,0.00069496455,0.0008337207,0.00027877736,0.0011147121,0.0009202474,0.000932956,0.003610791],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008075328,0.0005046601,0.017142883,0.0004216141,0.0001970619,0.00055185694,0.0011676435,0.009241265,0.0496588,0.0033976992,0.03681632,0.8800927],"study_design_scores_gemma":[0.00006199678,0.0005563797,0.06797236,0.00009665054,0.00021977736,0.0017080079,0.0012800936,0.83734995,0.04889409,0.009047187,0.03266096,0.00015245557],"about_ca_topic_score_codex":0.003358996,"about_ca_topic_score_gemma":0.006204543,"teacher_disagreement_score":0.0035938553,"about_ca_system_score_codex":0.00054785906,"about_ca_system_score_gemma":0.00042696914,"threshold_uncertainty_score":0.012022674},"labels":[],"label_agreement":null},{"id":"W2908510269","doi":"10.4000/books.aaccademia.4577","title":"Fully Convolutional Networks for Text Classification","year":2018,"lang":"en","type":"preprint","venue":"Accademia University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Research Council Canada; Università degli Studi di Napoli Federico II","keywords":"Emoji; Computer science; Task (project management); Convolutional neural network; Artificial intelligence; Machine learning; Data mining; World Wide Web; Social media; Engineering","score_opus":0.0656289605154812,"score_gpt":0.2576380090301897,"score_spread":0.19200904851470849,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2908510269","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021854835,0.006394219,0.94545704,0.0024417606,0.0008051749,0.00010326711,0.0026770653,0.0071298126,0.013136945],"genre_scores_gemma":[0.58855957,0.0056848596,0.33307767,0.0013565946,0.0012440758,0.00033636627,0.011215487,0.0008065591,0.057718713],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994143,0.00012295757,0.00003258996,0.00019450439,0.00013986025,0.00009570815],"domain_scores_gemma":[0.9991234,0.00039138712,0.00007553301,0.00017777627,0.00019469828,0.000037169804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070877525,0.0013878002,0.0005746095,0.0010468459,0.00041995844,0.0012196777,0.0013493913,0.0015291964,0.008385444],"category_scores_gemma":[0.0032459185,0.00048033477,0.0008248955,0.0014633023,0.00042322764,0.0027460286,0.00087997015,0.002007497,0.005144646],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030194697,0.00016925014,0.0015459439,0.00047896183,0.00024094073,0.00020303082,0.00015010269,0.16494234,0.021293588,0.04834472,0.053838883,0.70849025],"study_design_scores_gemma":[0.000011543468,0.000036693145,0.00084467384,0.00004920079,0.00003258754,0.000053175507,0.000017611676,0.9282354,0.0054371753,0.05193752,0.013320712,0.000023768633],"about_ca_topic_score_codex":0.010353685,"about_ca_topic_score_gemma":0.01581301,"teacher_disagreement_score":0.010353685,"about_ca_system_score_codex":0.0013024185,"about_ca_system_score_gemma":0.000884648,"threshold_uncertainty_score":0.028052151},"labels":[],"label_agreement":null},{"id":"W2909027478","doi":"10.1145/3284869.3284892","title":"Automated Extraction of Symptoms related to Rare Diseases from Scientific Publications","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Computer science; Exploit; Statistic; Focus (optics); Population; Task (project management); Term (time); Information extraction; Information retrieval; Data science; Medicine; Statistics; Computer security; Mathematics","score_opus":0.020345180790872438,"score_gpt":0.28706735857935817,"score_spread":0.2667221777884857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2909027478","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28875536,0.01705771,0.5718546,0.0040922128,0.0015183144,0.0012473481,0.06844911,0.03116332,0.015861982],"genre_scores_gemma":[0.37970412,0.007601781,0.5289427,0.00040230755,0.0013651032,0.00053352525,0.07554985,0.00050572056,0.0053948103],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99854004,0.00028409698,0.00033553076,0.0003986771,0.00035270725,0.00008892265],"domain_scores_gemma":[0.9912856,0.0045202267,0.0017258335,0.0006577782,0.0015750171,0.00023550517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014896741,0.0010852055,0.00090588664,0.013425823,0.0005564659,0.0018780751,0.00077438384,0.0011106653,0.0017439199],"category_scores_gemma":[0.008162438,0.000302253,0.001232699,0.0069143367,0.00028648236,0.002258955,0.0011235415,0.00062704546,0.0021654612],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035885477,0.00033376357,0.05441603,0.0032830292,0.00038713223,0.003391313,0.0012063764,0.0038073673,0.061253306,0.004479296,0.038949914,0.8281336],"study_design_scores_gemma":[0.00023851631,0.00078456366,0.2814781,0.0012305907,0.0019986348,0.014475979,0.0031319521,0.24308349,0.13674404,0.027929604,0.2884646,0.0004398662],"about_ca_topic_score_codex":0.0014884482,"about_ca_topic_score_gemma":0.0017557667,"teacher_disagreement_score":0.013425823,"about_ca_system_score_codex":0.00038583143,"about_ca_system_score_gemma":0.0016304564,"threshold_uncertainty_score":0.007878244},"labels":[],"label_agreement":null},{"id":"W2910577570","doi":"10.1145/3308774.3308781","title":"The Neural Hype and Comparisons Against Weak Baselines","year":2019,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":136,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Pace; Hyperparameter; Benchmark (surveying); Field (mathematics); Artificial intelligence; Data science; Point (geometry); Machine learning","score_opus":0.018584518233013887,"score_gpt":0.24437018413777653,"score_spread":0.22578566590476265,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2910577570","genre_codex":"review","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22230205,0.40950137,0.15835668,0.049268533,0.01877152,0.00057520944,0.009082236,0.007602289,0.124540135],"genre_scores_gemma":[0.86760426,0.020901384,0.06721763,0.005389227,0.004602121,0.0003324715,0.011856808,0.0019283141,0.020167757],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9818334,0.0096145235,0.00090107403,0.002605363,0.004448348,0.0005972666],"domain_scores_gemma":[0.9750175,0.014886116,0.0011121818,0.005420194,0.0027809413,0.0007830029],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026127232,0.0023517546,0.0019753585,0.0044492795,0.0020117671,0.004615481,0.0038883828,0.003661624,0.011257546],"category_scores_gemma":[0.07712321,0.00048361553,0.0010575816,0.0031008404,0.0030072138,0.010877495,0.005251765,0.006180594,0.006145902],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007188971,0.0008027356,0.010657815,0.0029031946,0.0030033025,0.00027338052,0.0004270255,0.051889688,0.002795865,0.055860367,0.15195742,0.7122403],"study_design_scores_gemma":[0.001074737,0.0062742964,0.024070302,0.0026945164,0.0023004215,0.0010924949,0.0013571337,0.5084532,0.011517838,0.2851199,0.15563886,0.00040626506],"about_ca_topic_score_codex":0.0044342284,"about_ca_topic_score_gemma":0.006640795,"teacher_disagreement_score":0.9738728,"about_ca_system_score_codex":0.0022284703,"about_ca_system_score_gemma":0.0010777449,"threshold_uncertainty_score":0.13817567},"labels":[],"label_agreement":null},{"id":"W2911227954","doi":"10.1162/tacl_a_00041","title":"Data Statements for Natural Language Processing: Toward Mitigating System Bias and Enabling Better Science","year":2018,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":808,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Macquarie University; York University; University of Washington; University of California, San Diego; National Science Foundation","keywords":"Embarrassment; Computer science; Natural (archaeology); Data science; Field (mathematics); Natural language; Lead (geology); Style (visual arts); Engineering ethics; Natural language processing; Psychology; Social psychology","score_opus":0.06724238139102369,"score_gpt":0.34862671363229547,"score_spread":0.2813843322412718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911227954","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068340334,0.00018225149,0.9646278,0.019342646,0.00040048416,0.00091623276,0.00038868672,0.0029390063,0.004368837],"genre_scores_gemma":[0.110289656,0.0002772553,0.87141496,0.00985027,0.00045422648,0.003254656,0.0010997661,0.0012828625,0.002076316],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.7011495,0.2327128,0.019739613,0.013440354,0.030838603,0.002119228],"domain_scores_gemma":[0.3909179,0.3800997,0.032558627,0.12905554,0.06091979,0.006448459],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.25439033,0.0012852662,0.0015191343,0.0039262646,0.0041392543,0.01144222,0.0053340346,0.004844823,0.005221077],"category_scores_gemma":[0.41203615,0.0020097855,0.0014945989,0.0037312629,0.013341458,0.03602486,0.016105805,0.012317075,0.0028403967],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009271836,0.00031649193,0.009593546,0.0022633388,0.0001574796,0.00045848347,0.017378602,0.0040050917,0.008293481,0.77331114,0.026017835,0.15727735],"study_design_scores_gemma":[0.0003561995,0.00047659667,0.0015266831,0.0015883291,0.00016884407,0.00047160243,0.0035564264,0.034604494,0.021619147,0.68600714,0.24937345,0.0002509995],"about_ca_topic_score_codex":0.0013721178,"about_ca_topic_score_gemma":0.00091061817,"teacher_disagreement_score":0.74560964,"about_ca_system_score_codex":0.0030595541,"about_ca_system_score_gemma":0.011024692,"threshold_uncertainty_score":0.9194695},"labels":[],"label_agreement":null},{"id":"W2911338114","doi":"10.1109/access.2019.2891692","title":"Challenging the Boundaries of Unsupervised Learning for Semantic Similarity","year":2019,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Semantic similarity; Artificial intelligence; Similarity (geometry); Natural language processing; Benchmark (surveying); Sentence; Word (group theory); Unsupervised learning; Mathematics","score_opus":0.04456077407634325,"score_gpt":0.29826225469932544,"score_spread":0.2537014806229822,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911338114","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033023704,0.0009364685,0.9614868,0.0012140914,0.00007834824,0.00016739001,0.00021171849,0.0005444471,0.0023370758],"genre_scores_gemma":[0.6270418,0.0007656063,0.36693755,0.0005741809,0.000530622,0.0005360475,0.0015182474,0.00040297664,0.0016929166],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98031163,0.011007127,0.0011446654,0.0044788886,0.0026820474,0.0003755496],"domain_scores_gemma":[0.9313058,0.0528606,0.0024792494,0.007298245,0.005120513,0.00093562773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019243238,0.0012006103,0.0022174888,0.0039106198,0.001990312,0.004541766,0.0036757016,0.0032291582,0.0010333532],"category_scores_gemma":[0.07397843,0.0006173057,0.001462181,0.0033183424,0.0034329304,0.009493399,0.0055730105,0.0045220843,0.00090043084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043335132,0.00081909203,0.015137317,0.0007408742,0.0006018046,0.00023279393,0.0012669328,0.16145015,0.0041906266,0.1623956,0.013958389,0.6387732],"study_design_scores_gemma":[0.000025106996,0.00008528871,0.0016943935,0.000054507836,0.000023722068,0.00011118983,0.00023782358,0.7916578,0.0016307042,0.20229006,0.002159198,0.000030227771],"about_ca_topic_score_codex":0.0027199166,"about_ca_topic_score_gemma":0.0027231004,"teacher_disagreement_score":0.019243238,"about_ca_system_score_codex":0.0019758982,"about_ca_system_score_gemma":0.002565922,"threshold_uncertainty_score":0.10176921},"labels":[],"label_agreement":null},{"id":"W2912277365","doi":"10.1109/bigdata.2018.8622463","title":"Deep Neural Networks for Social Media Word Segmentation of Asian Languages","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Word (group theory); Artificial intelligence; Natural language processing; Social media; Artificial neural network; Character (mathematics); Segmentation; Task (project management); Text segmentation; Language model; Layer (electronics); Linguistics; World Wide Web","score_opus":0.02608668151191412,"score_gpt":0.29246605207573306,"score_spread":0.26637937056381894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912277365","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4427762,0.004131419,0.5183073,0.0016294103,0.00043962518,0.00024702732,0.0043390566,0.0146910865,0.013438702],"genre_scores_gemma":[0.88868004,0.00086007774,0.092637286,0.00025930308,0.00014650624,0.00018726358,0.0059352606,0.00024940763,0.011044906],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997466,0.000058313468,0.000021918286,0.000085844425,0.000030976076,0.000056327142],"domain_scores_gemma":[0.9996822,0.00014475157,0.00005162147,0.000029169887,0.0000730467,0.000019283778],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00039774246,0.0012250884,0.00042563013,0.0013227334,0.0004931488,0.0006354231,0.0006078317,0.00065391086,0.002824005],"category_scores_gemma":[0.0012562566,0.0003335919,0.00062729657,0.0016005505,0.0003025987,0.0018243879,0.0007181837,0.0011712974,0.0016382057],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006790262,0.00035453797,0.006776321,0.0002810458,0.00020086643,0.00043066218,0.00064401847,0.17449538,0.024469784,0.0069617573,0.015625872,0.7690807],"study_design_scores_gemma":[0.000009943141,0.000028943381,0.0011697419,0.000013534037,0.000024254075,0.000024796873,0.00008759177,0.9871982,0.0044594957,0.004979046,0.0019933316,0.0000111117615],"about_ca_topic_score_codex":0.012866084,"about_ca_topic_score_gemma":0.017852057,"teacher_disagreement_score":0.012866084,"about_ca_system_score_codex":0.0007832794,"about_ca_system_score_gemma":0.0007442016,"threshold_uncertainty_score":0.025582373},"labels":[],"label_agreement":null},{"id":"W2912576237","doi":"10.1109/bibm.2018.8621452","title":"Recognising Named Entity of Medical Imaging Procedures in Clinical Notes","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Natural language processing; Conditional random field; Artificial intelligence; Medical imaging; Word (group theory); Process (computing); Named-entity recognition; Information retrieval; Programming language; Mathematics; Engineering","score_opus":0.0504459107749054,"score_gpt":0.376016117440868,"score_spread":0.3255702066659626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912576237","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33845398,0.010595574,0.49150395,0.0031632925,0.0013611459,0.0012495199,0.09331509,0.041736815,0.018620554],"genre_scores_gemma":[0.610266,0.002151372,0.24783818,0.0006419984,0.00047349912,0.00036036587,0.13026106,0.00058942474,0.007418098],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99772364,0.00043783762,0.00032753768,0.0008427429,0.0004793418,0.00018894898],"domain_scores_gemma":[0.99297225,0.0038591737,0.0010951455,0.00089904235,0.00091767625,0.00025655812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022509422,0.0009191043,0.00067576754,0.0055378973,0.00061055797,0.0013700448,0.0012333244,0.0019561185,0.0033088038],"category_scores_gemma":[0.008334592,0.00031019887,0.0013237882,0.0031557353,0.0003868254,0.0030531264,0.0010895405,0.001042522,0.0037875785],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017223355,0.00060830853,0.08445733,0.002498195,0.000364021,0.0030375917,0.0011219054,0.016877457,0.06986234,0.0054442715,0.084029965,0.7299763],"study_design_scores_gemma":[0.00021926546,0.0006985283,0.21956275,0.0010243908,0.0013136421,0.011595433,0.0016103403,0.34259766,0.19185106,0.015039724,0.21394983,0.000537387],"about_ca_topic_score_codex":0.0068570627,"about_ca_topic_score_gemma":0.008270491,"teacher_disagreement_score":0.0068570627,"about_ca_system_score_codex":0.0008677205,"about_ca_system_score_gemma":0.0015740943,"threshold_uncertainty_score":0.013634324},"labels":[],"label_agreement":null},{"id":"W2912817604","doi":"10.18653/v1/n19-4013","title":"End-to-End Open-Domain Question Answering with","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":368,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Question answering; Computer science; Open domain; Benchmark (surveying); Information retrieval; Reading (process); End-to-end principle; Natural language processing; Artificial intelligence; Domain (mathematical analysis); Reading comprehension; Language model; Open source; Programming language; Linguistics; Software","score_opus":0.028374056521567534,"score_gpt":0.2801999427192826,"score_spread":0.25182588619771507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912817604","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01717217,0.0016892639,0.7188481,0.0014060134,0.000807245,0.0007644944,0.014726072,0.23094389,0.013642687],"genre_scores_gemma":[0.18433848,0.0005287297,0.727059,0.0013322119,0.0004135293,0.0010719574,0.059789915,0.0051319385,0.020334255],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950918,0.0017034649,0.00033390938,0.0017308434,0.00080108555,0.00033889996],"domain_scores_gemma":[0.9934837,0.0033398953,0.0001484884,0.0018636374,0.00084654294,0.00031771767],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004099092,0.0024764426,0.0019339361,0.0023873341,0.0013461155,0.0039643557,0.0032767134,0.0033574225,0.026784804],"category_scores_gemma":[0.01264427,0.00096806494,0.001991447,0.0018461215,0.00079208944,0.0076367473,0.009914985,0.0031386812,0.034336247],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025435123,0.001018007,0.002410186,0.0011923249,0.00050674437,0.00093802076,0.0019023955,0.0077888127,0.020278923,0.018851645,0.35890865,0.5836607],"study_design_scores_gemma":[0.00050342793,0.00037504968,0.001926737,0.00021251822,0.00026147117,0.0006683469,0.001901749,0.5731444,0.044034924,0.1470115,0.22976667,0.00019319229],"about_ca_topic_score_codex":0.002672359,"about_ca_topic_score_gemma":0.0062928386,"teacher_disagreement_score":0.026784804,"about_ca_system_score_codex":0.0007565882,"about_ca_system_score_gemma":0.0010509599,"threshold_uncertainty_score":0.08960408},"labels":[],"label_agreement":null},{"id":"W2913282654","doi":"10.1145/3255771","title":"Session details: Research session 21: entity matching","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Session (web analytics); Computer science; Matching (statistics); World Wide Web; Statistics; Mathematics","score_opus":0.08264994715216234,"score_gpt":0.3641160850528131,"score_spread":0.2814661379006508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2913282654","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011255592,0.01822047,0.08265259,0.052407112,0.060804166,0.004451753,0.08683651,0.019855792,0.663516],"genre_scores_gemma":[0.043397367,0.00865542,0.021662086,0.0051257093,0.017825784,0.0017745596,0.075514436,0.003679018,0.8223657],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975061,0.00065486913,0.00013181627,0.00078226696,0.0006461235,0.00027879915],"domain_scores_gemma":[0.98796415,0.003561497,0.0002925524,0.002406146,0.0032014414,0.002574168],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0061467756,0.0020324613,0.0030601956,0.0017854646,0.0027463564,0.0068130433,0.001990834,0.0040004947,0.62643343],"category_scores_gemma":[0.01110196,0.0005257166,0.0022862973,0.0031156999,0.0004874683,0.005319078,0.0038548782,0.0030336787,0.51790565],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033243292,0.0002089755,0.0003503538,0.0003308329,0.000039792645,0.000049615683,0.00005316401,0.00018247413,0.002340261,0.0016254089,0.9275162,0.06697042],"study_design_scores_gemma":[0.00016263289,0.0003035069,0.002155333,0.00015972884,0.000107232365,0.0001666561,0.0001556349,0.0015804479,0.0031781888,0.0062400037,0.985748,0.00004258767],"about_ca_topic_score_codex":0.0019526567,"about_ca_topic_score_gemma":0.0033318282,"teacher_disagreement_score":0.37356657,"about_ca_system_score_codex":0.0012446435,"about_ca_system_score_gemma":0.0035920884,"threshold_uncertainty_score":0.53284734},"labels":[],"label_agreement":null},{"id":"W2914453913","doi":"10.1101/526244","title":"Towards reliable named entity recognition in the biomedical domain","year":2019,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Compute Canada; National Institutes of Health; Nvidia","keywords":"Conditional random field; CRFS; Overfitting; Computer science; Artificial intelligence; Dropout (neural networks); Named-entity recognition; Transfer of learning; Machine learning; Sequence labeling; Deep learning; Natural language processing; Task (project management); Multi-task learning; Regularization (linguistics); Generalization; Artificial neural network","score_opus":0.021337588013821525,"score_gpt":0.22800691114052032,"score_spread":0.2066693231266988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914453913","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041012716,0.008789764,0.8926106,0.006469633,0.0007762248,0.00027783963,0.010778983,0.03180419,0.007480014],"genre_scores_gemma":[0.28400773,0.0039142943,0.6536763,0.0018994836,0.0006395397,0.00031020743,0.04866145,0.0012026218,0.005688419],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935941,0.0025996312,0.00046371354,0.0016759888,0.001413271,0.00025331907],"domain_scores_gemma":[0.9781576,0.010577681,0.0016013137,0.0045860927,0.0045020026,0.0005753157],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01281725,0.0020373138,0.0014213406,0.0039435406,0.0008020824,0.0030892,0.002737934,0.0031701636,0.0048272563],"category_scores_gemma":[0.033421442,0.0006779753,0.0013220707,0.0035551146,0.0012879408,0.009203177,0.0043007885,0.003654302,0.011505194],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010669337,0.00039573005,0.01070845,0.0025094186,0.00033888672,0.0006994046,0.00063004106,0.1060251,0.037037365,0.029799003,0.122690715,0.688099],"study_design_scores_gemma":[0.0000743288,0.00021783306,0.00441912,0.00043514735,0.00014350799,0.000694094,0.0005072737,0.7966716,0.06739677,0.05918139,0.070129015,0.00012997152],"about_ca_topic_score_codex":0.0030938878,"about_ca_topic_score_gemma":0.0025107553,"teacher_disagreement_score":0.01281725,"about_ca_system_score_codex":0.0009224596,"about_ca_system_score_gemma":0.0023083724,"threshold_uncertainty_score":0.067784905},"labels":[],"label_agreement":null},{"id":"W2914713622","doi":"10.18653/v1/w18-5429","title":"Importance of Self-Attention for Sentiment Analysis","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"Natural Sciences and Engineering Research Council of Canada; Nvidia","keywords":"Interpretability; Computer science; Sentiment analysis; Task (project management); Artificial intelligence; Word (group theory); Baseline (sea); Machine learning; Natural language processing; Architecture; Order (exchange); Task analysis; Linguistics","score_opus":0.017395628524355753,"score_gpt":0.2703236454028812,"score_spread":0.2529280168785254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914713622","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19983426,0.0022910323,0.7836497,0.0018503453,0.00024408175,0.000102506834,0.00031530368,0.0029710373,0.008741785],"genre_scores_gemma":[0.9361156,0.00080879603,0.057604205,0.0002630405,0.0001502412,0.000052240688,0.00042423166,0.0001426959,0.0044389153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997092,0.00009990723,0.000014497022,0.0000885547,0.00005547634,0.000032283482],"domain_scores_gemma":[0.9988211,0.0007229668,0.000089492314,0.0001469033,0.00017227064,0.000047230926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009244239,0.00088928227,0.0003767034,0.0003632658,0.00020313855,0.0008259786,0.00061630394,0.0006734193,0.002598281],"category_scores_gemma":[0.0032586558,0.00022430727,0.00045595437,0.00033241027,0.00040587838,0.0019261845,0.0008020859,0.0012468285,0.00092176953],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044251376,0.00030111158,0.0091867875,0.00034219184,0.0002250411,0.00020366379,0.00039324712,0.1208091,0.06901467,0.01234455,0.0067665777,0.7799705],"study_design_scores_gemma":[0.000015993533,0.00016239347,0.0032974666,0.000037963324,0.00007475411,0.00008664264,0.000044884455,0.9564756,0.016124286,0.02034621,0.0033173885,0.000016381306],"about_ca_topic_score_codex":0.0015880262,"about_ca_topic_score_gemma":0.0031141262,"teacher_disagreement_score":0.002598281,"about_ca_system_score_codex":0.00041566612,"about_ca_system_score_gemma":0.00042224475,"threshold_uncertainty_score":0.008692145},"labels":[],"label_agreement":null},{"id":"W2915581179","doi":"10.48550/arxiv.1902.07249","title":"Discovery of Natural Language Concepts in Individual Units of CNNs","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Morpheme; Natural language processing; Artificial intelligence; Natural language; Natural (archaeology); Translation (biology); Machine translation; Linguistics; History","score_opus":0.07753615855421307,"score_gpt":0.21526531117439643,"score_spread":0.13772915262018337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2915581179","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.586745,0.0010972972,0.4033621,0.0010600879,0.00010969054,0.000106312764,0.0012067542,0.00085493055,0.0054577584],"genre_scores_gemma":[0.95468074,0.00025205838,0.0422403,0.00012906163,0.000036419686,0.00008578071,0.0009489691,0.00006747899,0.0015592827],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994875,0.000122005236,0.00002318568,0.00021854849,0.000063082036,0.00008562278],"domain_scores_gemma":[0.99872655,0.0006670611,0.00016859866,0.00018769297,0.00018043304,0.00006958775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008871678,0.0005628845,0.00040747892,0.000989346,0.00030809714,0.0010735032,0.00072900887,0.0006771184,0.001393799],"category_scores_gemma":[0.005046709,0.00040051676,0.0007474344,0.0009792816,0.0008292201,0.0025595753,0.00088702707,0.0014012714,0.00040882672],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010321351,0.00039661344,0.06088881,0.0008104691,0.00046993038,0.00065202254,0.0023268666,0.13868901,0.14504771,0.16874138,0.012563786,0.46838126],"study_design_scores_gemma":[0.000030924803,0.00008671223,0.013299686,0.000050419963,0.0000730746,0.00014203016,0.00026438385,0.85314995,0.017587945,0.11168948,0.003593961,0.000031412306],"about_ca_topic_score_codex":0.0031898513,"about_ca_topic_score_gemma":0.0033311443,"teacher_disagreement_score":0.0031898513,"about_ca_system_score_codex":0.0010094017,"about_ca_system_score_gemma":0.0006057748,"threshold_uncertainty_score":0.0073238015},"labels":[],"label_agreement":null},{"id":"W2916145857","doi":"10.1177/0165025419830248","title":"Using sentiment analysis to detect affect in children’s and adolescents’ poetry","year":2019,"lang":"en","type":"article","venue":"International Journal of Behavioral Development","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Sentiment analysis; Psychology; Valence (chemistry); Affect (linguistics); Context (archaeology); Developmental psychology; Social psychology; Natural language processing; Computer science; Chemistry; Communication","score_opus":0.024859029828473215,"score_gpt":0.3139432518610991,"score_spread":0.2890842220326259,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2916145857","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9875421,0.00017845853,0.007976716,0.00008918932,0.000030437915,0.00011850898,0.00088929554,0.000058517617,0.003116731],"genre_scores_gemma":[0.9824822,0.00024725826,0.015046921,0.00003381536,0.000030438132,0.00024436388,0.00094170956,0.000023989203,0.0009493132],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.999405,0.00019791594,0.000075562086,0.00011703005,0.00015974723,0.000044716955],"domain_scores_gemma":[0.9965976,0.0016931131,0.00077366,0.00014333877,0.00067172595,0.0001205953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001533576,0.00030398896,0.0002786429,0.0020009303,0.0003094911,0.0008878058,0.00015970317,0.00021568584,0.0016403409],"category_scores_gemma":[0.00691511,0.00017814335,0.00037835105,0.0012890538,0.0002349785,0.00076278916,0.00050289487,0.00044256638,0.00048154232],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003322629,0.00016675377,0.79342055,0.00027512407,0.00018072317,0.00029236698,0.00763315,0.0008635089,0.024196226,0.0009350738,0.0021958232,0.16950844],"study_design_scores_gemma":[0.000012873094,0.00018150693,0.97657436,0.000043590542,0.00006373723,0.00029643858,0.003892168,0.010557171,0.0040992997,0.0008547623,0.003400084,0.000024019962],"about_ca_topic_score_codex":0.0012246141,"about_ca_topic_score_gemma":0.0024639901,"teacher_disagreement_score":0.0020009303,"about_ca_system_score_codex":0.00029257147,"about_ca_system_score_gemma":0.0002270598,"threshold_uncertainty_score":0.008110464},"labels":[],"label_agreement":null},{"id":"W2916680872","doi":"","title":"THUNLP at TAC KBP 2013 in Entity Linking","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Task (project management); Computer science; Entity linking; Natural language processing; Artificial intelligence; F1 score; Knowledge base; Engineering","score_opus":0.009270023607219102,"score_gpt":0.2306121815772403,"score_spread":0.2213421579700212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2916680872","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09834956,0.005023728,0.27894527,0.0048051104,0.005552922,0.0029719784,0.08270019,0.43287426,0.08877708],"genre_scores_gemma":[0.21844853,0.0011717484,0.3329617,0.0020899433,0.0009179983,0.0022894617,0.34333402,0.02833096,0.07045566],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9881192,0.004535953,0.00086031586,0.0027245164,0.0028034002,0.00095668674],"domain_scores_gemma":[0.98791355,0.003271624,0.00031689293,0.0030668918,0.0043595973,0.0010714793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011861847,0.002487611,0.0018873187,0.004108128,0.0040540607,0.005197801,0.0033843978,0.002988572,0.031409133],"category_scores_gemma":[0.018022165,0.001330009,0.0012218922,0.0045467135,0.0008622836,0.010100913,0.0069967154,0.003806553,0.050549906],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018529943,0.0011117335,0.003979109,0.0012796962,0.00028065525,0.0011526507,0.0024974365,0.004172937,0.025680538,0.0034328867,0.61635035,0.33820903],"study_design_scores_gemma":[0.001045278,0.001491302,0.015224055,0.0003362318,0.00042214486,0.0024335124,0.0023794125,0.094836734,0.11725553,0.008091778,0.7559062,0.00057783397],"about_ca_topic_score_codex":0.022738578,"about_ca_topic_score_gemma":0.02392762,"teacher_disagreement_score":0.031409133,"about_ca_system_score_codex":0.0019037512,"about_ca_system_score_gemma":0.0028522143,"threshold_uncertainty_score":0.10507405},"labels":[],"label_agreement":null},{"id":"W2916853875","doi":"10.1007/978-3-030-11479-4_5","title":"Deep Learning for Document Representation","year":2019,"lang":"en","type":"book-chapter","venue":"Smart innovation, systems and technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ontario Institute of Technology","funders":"","keywords":"Word2vec; Computer science; Document classification; Representation (politics); Artificial intelligence; Natural language processing; Set (abstract data type); Information retrieval; Categorization; Bag-of-words model","score_opus":0.03384537849134838,"score_gpt":0.2627927763798293,"score_spread":0.2289473978884809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2916853875","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033623339,0.011036014,0.9600064,0.0013787652,0.0006522876,0.00006978434,0.0024859868,0.0058291717,0.015179137],"genre_scores_gemma":[0.11485001,0.01966969,0.70212346,0.0008190351,0.0011622655,0.0003588447,0.01783289,0.0015761598,0.14160764],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995989,0.00007516681,0.000032084907,0.00011904834,0.00013170225,0.000043143682],"domain_scores_gemma":[0.9993919,0.0002556455,0.00003217594,0.0001442928,0.0001518398,0.000024148405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00059635466,0.0009705395,0.0007832626,0.0014390806,0.0004106978,0.0017784178,0.0012521215,0.0010759307,0.013898537],"category_scores_gemma":[0.0019668317,0.00041661618,0.0007646038,0.002414693,0.00041952153,0.0029292447,0.0012276085,0.0020015342,0.010445811],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003561448,0.000038907143,0.00012695204,0.0002398343,0.000024182706,0.000028325749,0.000047307447,0.011854185,0.0042735967,0.022534918,0.0706696,0.8901266],"study_design_scores_gemma":[0.000016113605,0.000059587008,0.0008040093,0.00021626713,0.000058127906,0.00023326164,0.00007790846,0.59472,0.018246897,0.15973355,0.22578211,0.00005224758],"about_ca_topic_score_codex":0.0045522074,"about_ca_topic_score_gemma":0.0060353354,"teacher_disagreement_score":0.013898537,"about_ca_system_score_codex":0.0010511782,"about_ca_system_score_gemma":0.00083733135,"threshold_uncertainty_score":0.0464952},"labels":[],"label_agreement":null},{"id":"W2916976928","doi":"","title":"The TALP participation at TAC-KBP 2013","year":2013,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.011495903494673287,"score_gpt":0.24843168619303968,"score_spread":0.2369357826983664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2916976928","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06583893,0.005904795,0.42561203,0.06626286,0.029364677,0.007086301,0.08812147,0.052708343,0.25910056],"genre_scores_gemma":[0.1618559,0.0020367873,0.26147348,0.0074287183,0.004098066,0.005427924,0.18376347,0.02863361,0.34528193],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9828289,0.0062450506,0.00047000372,0.0024445874,0.006194392,0.0018169879],"domain_scores_gemma":[0.9715263,0.004497077,0.00037105277,0.0031365016,0.015198551,0.0052704853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02347066,0.0022857839,0.0021543803,0.0026872095,0.004996879,0.008669,0.004327467,0.0033579138,0.050677694],"category_scores_gemma":[0.031222397,0.0009651078,0.0016611083,0.0023765732,0.0018984943,0.0051907497,0.008552365,0.005798325,0.036044408],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012209155,0.0006254754,0.0007393971,0.0003953211,0.0001016091,0.0006636417,0.0029349306,0.0018534822,0.012685764,0.010630416,0.8683822,0.09976668],"study_design_scores_gemma":[0.00043070858,0.00023323143,0.0018522126,0.0002055863,0.00005923607,0.00031409273,0.0021446012,0.011329304,0.013368101,0.006204209,0.9637347,0.00012404406],"about_ca_topic_score_codex":0.047529638,"about_ca_topic_score_gemma":0.02493888,"teacher_disagreement_score":0.050677694,"about_ca_system_score_codex":0.0055470704,"about_ca_system_score_gemma":0.012626321,"threshold_uncertainty_score":0.16953373},"labels":[],"label_agreement":null},{"id":"W2917651062","doi":"10.1145/3303772.3303816","title":"Topic Development to Support Revision Feedback","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Paragraph; CLARITY; Computer science; Coherence (philosophical gambling strategy); Analytics; Meaning (existential); Data science; Psychology; World Wide Web","score_opus":0.025635896382171872,"score_gpt":0.25637911372882743,"score_spread":0.23074321734665557,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2917651062","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06550409,0.00039864815,0.8867146,0.0010286111,0.00057899754,0.0017767028,0.0013612531,0.03591106,0.006726075],"genre_scores_gemma":[0.27878746,0.00028729643,0.7030191,0.00030192887,0.0002401766,0.0027311346,0.002790127,0.0035892485,0.0082535045],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9911782,0.0044813147,0.0007281367,0.0017621001,0.0015919073,0.00025830237],"domain_scores_gemma":[0.91316533,0.051180452,0.0043049185,0.01147303,0.017548999,0.0023273046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011444808,0.0016606133,0.0009067605,0.0028388991,0.0009662915,0.004159243,0.00217119,0.0012432265,0.009865834],"category_scores_gemma":[0.09466945,0.000876958,0.0011874163,0.0014504278,0.00065083144,0.0042088176,0.004995238,0.0020234692,0.006424973],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010543049,0.0006390747,0.01228947,0.0017769081,0.00013137946,0.00048426227,0.022014167,0.009079286,0.04566591,0.00861022,0.03596175,0.8622933],"study_design_scores_gemma":[0.0008554288,0.0018096281,0.00943735,0.000985549,0.00042379723,0.0012825199,0.009265689,0.50358504,0.113570854,0.057567604,0.30077338,0.00044316173],"about_ca_topic_score_codex":0.0007482105,"about_ca_topic_score_gemma":0.001190411,"teacher_disagreement_score":0.011444808,"about_ca_system_score_codex":0.0009937519,"about_ca_system_score_gemma":0.002229899,"threshold_uncertainty_score":0.06052667},"labels":[],"label_agreement":null},{"id":"W2917934195","doi":"","title":"Cornell Belief and Sentiment System at TAC 2017.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence","score_opus":0.015634545266542843,"score_gpt":0.2464192913199238,"score_spread":0.23078474605338095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2917934195","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2192199,0.0014667495,0.05140059,0.0074646086,0.0028934819,0.0012206138,0.51854306,0.0918974,0.105893604],"genre_scores_gemma":[0.33873487,0.00027685356,0.048167832,0.00064405624,0.00045743032,0.0007304583,0.5712165,0.0022275213,0.037544534],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987531,0.00039982228,0.000062800595,0.00022107684,0.00042366213,0.00013968017],"domain_scores_gemma":[0.99668497,0.00083493546,0.00010942214,0.00055734866,0.0014827097,0.00033059865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031215542,0.0008851379,0.000562969,0.0022136276,0.0010658675,0.0018095347,0.0011731063,0.0010875708,0.017365199],"category_scores_gemma":[0.011488378,0.00031190907,0.00047991303,0.0014970674,0.0002887201,0.0029081742,0.0012531959,0.0016361901,0.013589447],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057545915,0.00037660613,0.009656941,0.00016365915,0.00009153529,0.00010069285,0.00026524224,0.0047464543,0.0011696443,0.0026454364,0.90932345,0.0708849],"study_design_scores_gemma":[0.0005030257,0.0003136164,0.033469908,0.00014965169,0.00019487905,0.00011889598,0.0006535607,0.48266512,0.009730711,0.0159712,0.45603484,0.00019455519],"about_ca_topic_score_codex":0.058348943,"about_ca_topic_score_gemma":0.1089298,"teacher_disagreement_score":0.058348943,"about_ca_system_score_codex":0.0022910873,"about_ca_system_score_gemma":0.0020357575,"threshold_uncertainty_score":0.11601865},"labels":[],"label_agreement":null},{"id":"W2918138400","doi":"10.18653/v1/2024.textgraphs-1","title":"Proceedings of TextGraphs-17: Graph-based Methods for Natural Language Processing","year":2024,"lang":"en","type":"paratext","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Division of Computer and Network Systems; Moscow Institute of Physics and Technology; Nazarbayev University; Bundesministerium für Bildung und Forschung; Natural Sciences and Engineering Research Council of Canada; European Commission; Analytical Center for the Government of the Russian Federation; University of Michigan; Automotive Research Center; National Science Foundation","keywords":"Computer science; Graph; Programming language; Natural language processing; Artificial intelligence; Theoretical computer science","score_opus":0.02670105931340308,"score_gpt":0.36625907457166823,"score_spread":0.3395580152582651,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2918138400","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030845352,0.024155049,0.9303745,0.0055452366,0.0065238923,0.00030765295,0.0030601446,0.006522474,0.020426545],"genre_scores_gemma":[0.030924564,0.02976078,0.8413957,0.0024158307,0.0062190485,0.0007375685,0.017440956,0.007196862,0.06390876],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966703,0.0012501825,0.00021088262,0.0008095467,0.00095366454,0.000105410596],"domain_scores_gemma":[0.99107856,0.0053658877,0.00014205292,0.0014403255,0.0015086106,0.00046455694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005029938,0.0021655622,0.0019465504,0.004413651,0.001137307,0.0038004033,0.0027454959,0.001755673,0.022262737],"category_scores_gemma":[0.011033272,0.0011308495,0.0016846386,0.0054348274,0.0014712617,0.005560868,0.0025967169,0.005519923,0.011641725],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016211711,0.00020962335,0.00039196128,0.00079657015,0.00027265024,0.000115851704,0.00043808162,0.009305236,0.0028754754,0.06330942,0.3321311,0.58999187],"study_design_scores_gemma":[0.000076158794,0.000075860575,0.0007262478,0.00026156814,0.000108661734,0.00021218606,0.00010771201,0.05091948,0.0062497337,0.13472246,0.80646974,0.00007023735],"about_ca_topic_score_codex":0.0075172163,"about_ca_topic_score_gemma":0.009456636,"teacher_disagreement_score":0.022262737,"about_ca_system_score_codex":0.002010501,"about_ca_system_score_gemma":0.0025819293,"threshold_uncertainty_score":0.07447624},"labels":[],"label_agreement":null},{"id":"W2922857","doi":"10.1007/978-3-642-10238-7_7","title":"Artificial K-Lines and Applications","year":2009,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Artificial intelligence; Artificial neural network; Causality (physics); Knowledge base; Machine learning; Physics","score_opus":0.06177799720094484,"score_gpt":0.3040433938657268,"score_spread":0.24226539666478197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2922857","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02454005,0.015757034,0.6594457,0.0031375447,0.0008564038,0.00007480105,0.00036927307,0.0016313483,0.29418793],"genre_scores_gemma":[0.48543113,0.016138086,0.3445472,0.0006337627,0.00074578775,0.00021561065,0.00073206145,0.0008016096,0.1507548],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99961877,0.00012669426,0.000030827923,0.00008806867,0.00010581675,0.00002978343],"domain_scores_gemma":[0.99874663,0.0006591227,0.00008896616,0.00020189611,0.00024031589,0.00006316942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005485195,0.0004102945,0.00046492004,0.0010081524,0.00094887894,0.0027175082,0.00088533486,0.00093402277,0.015502555],"category_scores_gemma":[0.0035237987,0.00030690889,0.00038208568,0.002217035,0.0015004455,0.0042768396,0.0010464047,0.0014649698,0.0055312733],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000036021123,0.00002223648,0.00024579847,0.00019530322,0.000010646747,0.000046377034,0.00027138967,0.0097060045,0.00082038215,0.8247568,0.014825735,0.1490633],"study_design_scores_gemma":[0.000008847534,0.000018720695,0.00013056079,0.0000834408,0.0000075662333,0.00016364547,0.00013490605,0.03641924,0.0007443614,0.83909154,0.12318209,0.000015044567],"about_ca_topic_score_codex":0.00060266134,"about_ca_topic_score_gemma":0.0005373483,"teacher_disagreement_score":0.015502555,"about_ca_system_score_codex":0.0009103527,"about_ca_system_score_gemma":0.0004751948,"threshold_uncertainty_score":0.051861167},"labels":[],"label_agreement":null},{"id":"W2923261428","doi":"10.48550/arxiv.1903.10126","title":"Connecting Language and Knowledge with Heterogeneous Representations for Neural Relation Extraction","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Task (project management); Computer science; Relation (database); Code (set theory); Natural language processing; Relationship extraction; Artificial intelligence; Constant (computer programming); Data science; Programming language; Data mining; Engineering","score_opus":0.0808915080268491,"score_gpt":0.23691885154855963,"score_spread":0.15602734352171055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2923261428","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03778127,0.0021124799,0.94799656,0.0020403222,0.00018369277,0.000093997245,0.0015974207,0.0035459392,0.004648404],"genre_scores_gemma":[0.6737883,0.0021521398,0.30237478,0.0009635485,0.00045950094,0.00040203065,0.0072318986,0.00058000395,0.012047774],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999246,0.00025847915,0.000050040995,0.00027731084,0.0000939103,0.000074378346],"domain_scores_gemma":[0.9976667,0.0014571866,0.00015057565,0.00044945115,0.00019376712,0.0000823771],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012219354,0.0010565439,0.0007323795,0.0022762874,0.00057953136,0.0019218654,0.0021845659,0.0018513942,0.0065619554],"category_scores_gemma":[0.008699056,0.000816954,0.0017332227,0.0031820661,0.0007680785,0.0066810376,0.0026296459,0.0031004897,0.0026985828],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000313634,0.0003175696,0.0039314646,0.000394236,0.0003630617,0.00034886628,0.00060125283,0.21568172,0.00617873,0.064903855,0.022952083,0.68401355],"study_design_scores_gemma":[0.000016885708,0.000029500103,0.0004572797,0.00005705621,0.00005451517,0.00008626206,0.00007821317,0.8614994,0.0015924467,0.13212466,0.003983092,0.000020630954],"about_ca_topic_score_codex":0.0054591917,"about_ca_topic_score_gemma":0.0112720765,"teacher_disagreement_score":0.0065619554,"about_ca_system_score_codex":0.0012451006,"about_ca_system_score_gemma":0.00079104595,"threshold_uncertainty_score":0.021951973},"labels":[],"label_agreement":null},{"id":"W2923796912","doi":"10.3233/icg-190094","title":"Hex 2018: MoHex3HNN over DeepEzo","year":2019,"lang":"en","type":"article","venue":"ICGA Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science","score_opus":0.014552622009464231,"score_gpt":0.23872388482465218,"score_spread":0.22417126281518795,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2923796912","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023098303,0.015574001,0.29451483,0.07682637,0.07251542,0.0006092611,0.06465493,0.080427304,0.37177953],"genre_scores_gemma":[0.11452837,0.0051282137,0.099430285,0.00462198,0.0068521355,0.00033068724,0.06604626,0.022606919,0.6804551],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99878556,0.00026391822,0.000052356398,0.00034327485,0.00037141985,0.00018356025],"domain_scores_gemma":[0.9976363,0.0004279676,0.00007459283,0.00058376114,0.00059047475,0.0006868738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033122825,0.0010575933,0.0011995997,0.00089541933,0.0013724563,0.004118581,0.0012301726,0.0011322227,0.21734704],"category_scores_gemma":[0.006389795,0.00038173687,0.000596896,0.000914625,0.00063714996,0.0031587097,0.004349519,0.0017007403,0.097467884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005507941,0.00006793811,0.00039323673,0.00021904628,0.000034327335,0.00009730786,0.00013139173,0.0030685372,0.0019123803,0.031870563,0.7436779,0.21797654],"study_design_scores_gemma":[0.00015223949,0.000058820977,0.000514483,0.00016028053,0.00001733172,0.00009265575,0.000103449165,0.021587912,0.0020478377,0.05100957,0.92422336,0.000032062067],"about_ca_topic_score_codex":0.0091696,"about_ca_topic_score_gemma":0.012545699,"teacher_disagreement_score":0.21734704,"about_ca_system_score_codex":0.0016657036,"about_ca_system_score_gemma":0.0031207292,"threshold_uncertainty_score":0.7270983},"labels":[],"label_agreement":null},{"id":"W2923890923","doi":"10.48550/arxiv.1903.10972","title":"Models and Data for Simple Applications of BERT for Ad Hoc Document Retrieval","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":133,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Simple (philosophy); Sentence; Information retrieval; Microblogging; Inference; Question answering; Social media; Post hoc; Artificial intelligence; Natural language processing; World Wide Web","score_opus":0.15109454832975105,"score_gpt":0.24884759872786272,"score_spread":0.09775305039811166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2923890923","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044754636,0.0011626633,0.9284828,0.002768616,0.00017824305,0.00035126804,0.0076560415,0.004924961,0.009720769],"genre_scores_gemma":[0.5416336,0.001171716,0.4225236,0.0006687716,0.00059183757,0.0011356613,0.02071628,0.00058235973,0.010976268],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99871147,0.0005011586,0.00010445364,0.0002761238,0.00031178954,0.0000949815],"domain_scores_gemma":[0.9931538,0.004322365,0.0004321149,0.0012761487,0.00067794626,0.00013769072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004184459,0.00090972794,0.0007430246,0.0018738483,0.0007886645,0.0019551904,0.002555566,0.0017233061,0.00855207],"category_scores_gemma":[0.01878857,0.0006361115,0.00092670793,0.002249957,0.00072004023,0.004319523,0.0010872512,0.0024760398,0.0037510507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075970095,0.00041632887,0.011452266,0.0005800502,0.0001588371,0.00038972564,0.00051860866,0.51451474,0.005382609,0.20424996,0.04310464,0.2184725],"study_design_scores_gemma":[0.000026833899,0.000040799834,0.0013144707,0.000027062522,0.000014123947,0.00012314199,0.000042899672,0.9096616,0.0010788372,0.07751178,0.010129514,0.000028856603],"about_ca_topic_score_codex":0.011272205,"about_ca_topic_score_gemma":0.021266941,"teacher_disagreement_score":0.011272205,"about_ca_system_score_codex":0.0019925816,"about_ca_system_score_gemma":0.0011959614,"threshold_uncertainty_score":0.028609514},"labels":[],"label_agreement":null},{"id":"W2928075308","doi":"10.21437/interspeech.2019-2396","title":"Speech Model Pre-Training for End-to-End Spoken Language Understanding","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; McGill University","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Canadian Institute for Advanced Research","keywords":"Computer science; End-to-end principle; Training set; Spoken language; Speech recognition; Training (meteorology); Artificial intelligence; Language model; Natural language processing; Gauge (firearms); Speech synthesis","score_opus":0.1407860486826978,"score_gpt":0.33036668898392463,"score_spread":0.18958064030122682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2928075308","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04487442,0.00064240413,0.91350174,0.00033175244,0.00023379149,0.00023789385,0.0031212938,0.033882927,0.003173731],"genre_scores_gemma":[0.47546285,0.00042076773,0.4895579,0.00048295004,0.00014816181,0.0011305403,0.021179583,0.0020649952,0.0095523],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99892575,0.0003420095,0.000068075344,0.00040470302,0.00015191552,0.00010751036],"domain_scores_gemma":[0.99673235,0.0016869262,0.00008727845,0.0007663233,0.0006204804,0.000106639025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013439391,0.0018304407,0.0011354316,0.00066625513,0.00061817403,0.0011494406,0.001792905,0.0016630009,0.007618682],"category_scores_gemma":[0.0062288614,0.00069050066,0.0011390428,0.00058191054,0.0004364251,0.0028467688,0.0017007103,0.0037507014,0.0076545984],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074991764,0.0006322336,0.003983825,0.0005195258,0.0002966038,0.0003206223,0.0009350733,0.14193176,0.06253775,0.0040293536,0.04890985,0.73515356],"study_design_scores_gemma":[0.000040088857,0.00016784812,0.0013735838,0.000029528172,0.00005018994,0.00011931662,0.00020599533,0.94306874,0.042324666,0.00437691,0.008200499,0.000042666965],"about_ca_topic_score_codex":0.006913005,"about_ca_topic_score_gemma":0.014862965,"teacher_disagreement_score":0.007618682,"about_ca_system_score_codex":0.00069305865,"about_ca_system_score_gemma":0.0015339932,"threshold_uncertainty_score":0.025487065},"labels":[],"label_agreement":null},{"id":"W2929767294","doi":"10.18653/v1/n19-1381","title":"Evaluating Coherence in Dialogue Systems using Entailment","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Coherence (philosophical gambling strategy); Textual entailment; Sentence; Natural language processing; Artificial intelligence; Logical consequence; Scalability; Quality (philosophy); Meaning (existential); Machine learning; Mathematics; Statistics; Psychology; Epistemology","score_opus":0.20563514454950158,"score_gpt":0.37735979186949564,"score_spread":0.17172464731999407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2929767294","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55922985,0.0066715893,0.41773272,0.0016429067,0.0002499137,0.0007520326,0.0017851195,0.0024205772,0.009515191],"genre_scores_gemma":[0.93170565,0.0003578281,0.064675614,0.000054908014,0.00008756211,0.00017436456,0.0022194218,0.0001244182,0.0006003049],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9842626,0.010173927,0.0009885082,0.0017727464,0.0023525825,0.00044966032],"domain_scores_gemma":[0.96412104,0.028510323,0.0017985534,0.0017647573,0.0032252022,0.0005801031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077385544,0.0008052575,0.0010565303,0.0029686973,0.001475936,0.0035561197,0.0009932333,0.0016740622,0.0025567426],"category_scores_gemma":[0.05157032,0.000623418,0.00074544235,0.001521712,0.0010390591,0.0062872977,0.0043085855,0.0013384834,0.00078521157],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008018029,0.0009941729,0.040765427,0.0037792886,0.0016747965,0.00094363804,0.010000319,0.0959182,0.045540277,0.057119697,0.0123136295,0.7229325],"study_design_scores_gemma":[0.00032210202,0.0011483297,0.014423811,0.00033421008,0.000606624,0.0003515216,0.0031091482,0.8501715,0.028538272,0.09156652,0.009269215,0.00015870435],"about_ca_topic_score_codex":0.0032375355,"about_ca_topic_score_gemma":0.0039985497,"teacher_disagreement_score":0.0077385544,"about_ca_system_score_codex":0.0012327313,"about_ca_system_score_gemma":0.001377772,"threshold_uncertainty_score":0.04092592},"labels":[],"label_agreement":null},{"id":"W2935885634","doi":"10.1145/3331184.3331296","title":"An Axiomatic Approach to Regularizing Neural Ranking Models","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Axiom; Axiomatic system; Ranking (information retrieval); Regularization (linguistics); Machine learning; Computer science; Generalization; Artificial neural network; Artificial intelligence; Relevance (law); Convergence (economics); Training set; Set (abstract data type); Mathematical optimization; Algorithm; Mathematics","score_opus":0.06637320565648477,"score_gpt":0.27037379469780237,"score_spread":0.2040005890413176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2935885634","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045779394,0.00010571276,0.9922126,0.00046800348,0.000035262547,0.000022832855,0.00008625244,0.00024217863,0.0022492507],"genre_scores_gemma":[0.30617544,0.0004581948,0.6818885,0.0008676931,0.0003188839,0.0004068784,0.00073887844,0.0003074299,0.008838074],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9955556,0.0023593714,0.00027579858,0.00070277817,0.00095175335,0.0001546497],"domain_scores_gemma":[0.9895579,0.005634049,0.00081903266,0.0025249296,0.0012743878,0.00018970565],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006202021,0.00083568215,0.00092560524,0.0011934214,0.00075135217,0.0018030548,0.0026245282,0.0017380806,0.003576409],"category_scores_gemma":[0.025211083,0.00064524263,0.0014081675,0.0014313345,0.002321353,0.005010265,0.0023317137,0.0048939763,0.0009996197],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006172729,0.000085403444,0.0007296126,0.00013865216,0.00007455311,0.0000806782,0.0002214316,0.20347288,0.0020504051,0.7342865,0.0049162954,0.053881753],"study_design_scores_gemma":[0.000012738482,0.000036128473,0.0001458433,0.000019734976,0.000014063458,0.00005056415,0.000016678423,0.6090588,0.0006781144,0.38665298,0.0032945394,0.000019814996],"about_ca_topic_score_codex":0.0022944075,"about_ca_topic_score_gemma":0.0041919164,"teacher_disagreement_score":0.006202021,"about_ca_system_score_codex":0.0016399417,"about_ca_system_score_gemma":0.0015178495,"threshold_uncertainty_score":0.03279984},"labels":[],"label_agreement":null},{"id":"W2940263831","doi":"10.48550/arxiv.1904.09171","title":"Critically Examining the \"Neural Hype\": Weak Baselines and the Additivity of Effectiveness Gains from Neural Ranking Models","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Skepticism; Ranking (information retrieval); Computer science; Artificial neural network; Post hoc; Artificial intelligence; Additive function; Deep neural networks; Machine learning; Epistemology; Mathematics; Philosophy; Medicine","score_opus":0.10987573863177681,"score_gpt":0.20786419278357626,"score_spread":0.09798845415179945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2940263831","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17035216,0.22800922,0.40158796,0.088645905,0.0045125694,0.00079253863,0.0037075975,0.0032201817,0.09917193],"genre_scores_gemma":[0.88100225,0.01095435,0.091150805,0.007265763,0.0027573884,0.00038513372,0.0013326899,0.0005552337,0.0045964075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94775146,0.038729902,0.0020456244,0.0030618669,0.0078772195,0.00053392374],"domain_scores_gemma":[0.82474965,0.14729138,0.0041489284,0.014718729,0.008228357,0.0008630356],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.075879976,0.0020613188,0.002581986,0.0040646945,0.0011193979,0.005175572,0.0029770345,0.0019851206,0.0040177736],"category_scores_gemma":[0.17182568,0.00070652744,0.0015045883,0.002994447,0.0033182683,0.013270342,0.0032097981,0.0066226986,0.0017516696],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0041643805,0.0006105772,0.013915624,0.0069061257,0.009943246,0.00016056746,0.00074564037,0.06241634,0.00404087,0.12240951,0.041757148,0.73293],"study_design_scores_gemma":[0.0015043668,0.0057638222,0.017470602,0.0036604484,0.008852546,0.00051247043,0.0009371135,0.3451596,0.017301852,0.5383293,0.05995678,0.0005510719],"about_ca_topic_score_codex":0.0022705733,"about_ca_topic_score_gemma":0.004341376,"teacher_disagreement_score":0.92412,"about_ca_system_score_codex":0.0019631605,"about_ca_system_score_gemma":0.0016012532,"threshold_uncertainty_score":0.40129644},"labels":[],"label_agreement":null},{"id":"W2940730617","doi":"10.2196/11499","title":"Adapting State-of-the-Art Deep Language Models to Clinical Information Extraction Systems: Potentials, Challenges, and Solutions","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Department of Education and Training; Australian National University","keywords":"Computer science; Artificial intelligence; Natural language processing; Information extraction; Vocabulary; Context (archaeology); Health informatics; Informatics; Natural language; Machine learning; Data science; Health care; Human–computer interaction","score_opus":0.052890566908270424,"score_gpt":0.3127531179032773,"score_spread":0.2598625509950069,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2940730617","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09093706,0.012290781,0.87237936,0.012723348,0.00046369637,0.000290881,0.0005060402,0.005315078,0.005093704],"genre_scores_gemma":[0.58778125,0.007847837,0.39532086,0.0024594956,0.00035988793,0.00035809056,0.0015124836,0.0003125866,0.00404745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978497,0.0010307429,0.00019160674,0.00041871003,0.00034225712,0.00016698471],"domain_scores_gemma":[0.9938122,0.003825539,0.00026169978,0.0009058928,0.0010310763,0.00016363218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006103388,0.0011380895,0.00068723946,0.0010364199,0.00031922976,0.0020998698,0.0025188746,0.0019096934,0.00193195],"category_scores_gemma":[0.013553049,0.0006036993,0.0011002,0.001118256,0.0010423596,0.0040406617,0.0021862972,0.0026773482,0.00122064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002913504,0.0004938957,0.006040737,0.0006973973,0.00024744237,0.00025947148,0.00053878815,0.19645365,0.014117424,0.008515953,0.007736362,0.7646075],"study_design_scores_gemma":[0.000028074906,0.00018771009,0.00087256974,0.00011496942,0.00005721074,0.00010246346,0.00019732199,0.9629373,0.011228315,0.016892467,0.007340917,0.000040687944],"about_ca_topic_score_codex":0.0062639723,"about_ca_topic_score_gemma":0.005225802,"teacher_disagreement_score":0.0062639723,"about_ca_system_score_codex":0.0014649761,"about_ca_system_score_gemma":0.001874628,"threshold_uncertainty_score":0.03227818},"labels":[],"label_agreement":null},{"id":"W2941044695","doi":"10.1109/besc.2018.8697314","title":"Latent Semantic Analysis Boosted Convolutional Neural Networks for Document Classification","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"tf–idf; Computer science; Convolutional neural network; Artificial intelligence; Latent semantic analysis; Weighting; Trigram; Singular value decomposition; Word (group theory); Pattern recognition (psychology); Word2vec; Support vector machine; Machine learning; Natural language processing; Term (time); Mathematics","score_opus":0.041723874394951656,"score_gpt":0.2814619927427373,"score_spread":0.23973811834778566,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2941044695","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10128328,0.007877023,0.870481,0.0009663517,0.00042371306,0.00013739547,0.0029726187,0.009280322,0.0065782834],"genre_scores_gemma":[0.70322907,0.003961165,0.26945558,0.00035104653,0.00029086132,0.0001662662,0.007415373,0.0002719521,0.014858735],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99964297,0.000054826804,0.00002698437,0.00011546169,0.00009717176,0.000062693536],"domain_scores_gemma":[0.9995772,0.00014522388,0.00006811603,0.000068978865,0.000120511075,0.000019894656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005599975,0.0009813448,0.0006261763,0.0011167055,0.00027579666,0.000791415,0.0009044937,0.0006155393,0.0031063582],"category_scores_gemma":[0.0015067531,0.00030506027,0.0008809209,0.001532385,0.00029952466,0.0014448074,0.0005646373,0.0013748248,0.0017451323],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000319627,0.00023335578,0.0031577838,0.00036232994,0.00022752723,0.00013600687,0.00009065611,0.13971274,0.024904264,0.010811873,0.014647672,0.80539626],"study_design_scores_gemma":[0.000007693681,0.00003870521,0.0009284658,0.000021123502,0.00003383386,0.00003478255,0.000011760934,0.98479605,0.0054303682,0.006129216,0.0025535955,0.000014389338],"about_ca_topic_score_codex":0.012804881,"about_ca_topic_score_gemma":0.018211514,"teacher_disagreement_score":0.012804881,"about_ca_system_score_codex":0.0013259818,"about_ca_system_score_gemma":0.0009596028,"threshold_uncertainty_score":0.02546066},"labels":[],"label_agreement":null},{"id":"W2941057185","doi":"10.3389/fpsyg.2019.00825","title":"Multiple-Choice Item Distractor Development Using Topic Modeling Approaches","year":2019,"lang":"en","type":"article","venue":"Frontiers in Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Psychology; Task (project management); Multiple choice; Latent Dirichlet allocation; Test (biology); Quality (philosophy); Process (computing); Cognitive psychology; Natural language processing; Item response theory; Artificial intelligence; Topic model; Computer science; Mathematics education; Linguistics; Psychometrics; Developmental psychology","score_opus":0.09783407409281254,"score_gpt":0.2976876296197643,"score_spread":0.19985355552695178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2941057185","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027570851,0.0001821078,0.9667995,0.00014659046,0.00007570186,0.001814096,0.00028283685,0.0022499436,0.0008784644],"genre_scores_gemma":[0.08512813,0.00010223066,0.90745836,0.00010889592,0.000038059785,0.005003817,0.0006587864,0.00048284954,0.0010188315],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.93970233,0.040101558,0.0050889743,0.0073282206,0.00725969,0.0005191954],"domain_scores_gemma":[0.58534515,0.35956165,0.009770807,0.01879034,0.025323384,0.0012086359],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.055735357,0.0029209114,0.0019668122,0.0048407605,0.0016607511,0.0045511187,0.0046942397,0.0024591135,0.0060292855],"category_scores_gemma":[0.2027386,0.0020677447,0.0022711158,0.003486297,0.0014442386,0.0037204393,0.0053497273,0.0035602616,0.0025614237],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015917666,0.0010102862,0.025272317,0.002602308,0.00072912325,0.0005816524,0.015376367,0.015481205,0.03857702,0.014240623,0.0066233696,0.877914],"study_design_scores_gemma":[0.0011969659,0.0021841829,0.027811036,0.0010806175,0.0010970681,0.0021528534,0.0061048022,0.6920572,0.13901514,0.0706757,0.055852156,0.0007723625],"about_ca_topic_score_codex":0.0009786782,"about_ca_topic_score_gemma":0.0024495032,"teacher_disagreement_score":0.055735357,"about_ca_system_score_codex":0.0018555429,"about_ca_system_score_gemma":0.0025724543,"threshold_uncertainty_score":0.2947603},"labels":[],"label_agreement":null},{"id":"W2942156930","doi":"10.1109/access.2019.2910145","title":"Automating Articulation: Applying Natural Language Processing to Post-Secondary Credit Transfer","year":2019,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Transfer of learning; Parsing; Dependency grammar; Task (project management); Domain (mathematical analysis); Word2vec; Field (mathematics); Dependency (UML); Articulation (sociology); Subject-matter expert; Machine learning; Expert system","score_opus":0.015865662365712917,"score_gpt":0.28713390569065017,"score_spread":0.27126824332493726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2942156930","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.344387,0.00019218348,0.63834333,0.0003435005,0.00011227958,0.000642849,0.0011675202,0.0065758405,0.008235379],"genre_scores_gemma":[0.7550863,0.00009724723,0.23967029,0.000055286917,0.00006792351,0.00032596025,0.0021771993,0.0002214083,0.0022983542],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967002,0.0015844245,0.00024247858,0.00067169813,0.00061089464,0.00019022072],"domain_scores_gemma":[0.99034834,0.006070536,0.0009357816,0.0007280365,0.0017337984,0.00018349297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032655492,0.00061407645,0.00037031723,0.0035411082,0.0007976776,0.0019971621,0.0007487676,0.0006423934,0.0019274197],"category_scores_gemma":[0.011453547,0.00022458697,0.0005791754,0.0016798706,0.0005865135,0.0015948291,0.0017327417,0.00096167775,0.0016349939],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047366956,0.00049655157,0.024913771,0.0003275995,0.00007913699,0.00064240344,0.0051807873,0.025251484,0.04994284,0.0039119856,0.0042242417,0.8845555],"study_design_scores_gemma":[0.000073713534,0.0005982056,0.07100282,0.00011529489,0.00010221587,0.00071190594,0.006736456,0.76824874,0.098370716,0.027420936,0.0264511,0.00016783981],"about_ca_topic_score_codex":0.003542532,"about_ca_topic_score_gemma":0.0034170875,"teacher_disagreement_score":0.003542532,"about_ca_system_score_codex":0.0008168233,"about_ca_system_score_gemma":0.0014378724,"threshold_uncertainty_score":0.017270029},"labels":[],"label_agreement":null},{"id":"W2944617758","doi":"10.1111/cogs.12730","title":"The Role of Negative Information in Distributional Semantic Learning","year":2019,"lang":"en","type":"article","venue":"Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Word2vec; Computer science; Word embedding; Artificial intelligence; Context (archaeology); Natural language processing; Word (group theory); Representation (politics); Semantics (computer science); Machine learning; Embedding; Mathematics","score_opus":0.007649209568656575,"score_gpt":0.23724899860659293,"score_spread":0.22959978903793635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2944617758","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10964022,0.0026061262,0.86633444,0.005074501,0.00038298313,0.00009025839,0.0005121192,0.00084233395,0.014516986],"genre_scores_gemma":[0.9306716,0.001143092,0.061122943,0.001352901,0.00050194917,0.00011914734,0.00090976775,0.00029965126,0.003879116],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948783,0.0021590185,0.00028871553,0.001407314,0.001100385,0.00016632274],"domain_scores_gemma":[0.9708076,0.02159902,0.001436756,0.0031017272,0.0023300513,0.00072476536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008051501,0.0015955385,0.001518978,0.001880542,0.0012078201,0.0036636093,0.0025556656,0.0019703647,0.0034053025],"category_scores_gemma":[0.051912628,0.0008956694,0.0009233917,0.0014860696,0.005814232,0.014615161,0.0047306656,0.0043138154,0.00067619304],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00109297,0.0003563083,0.023414776,0.0006188631,0.0002711443,0.00053800305,0.0017555053,0.06187776,0.0070780166,0.55288,0.009073834,0.34104285],"study_design_scores_gemma":[0.00003417747,0.00009910367,0.0017928474,0.00007091062,0.000035778587,0.00026068714,0.00010239222,0.16432807,0.0018651538,0.8282825,0.003076851,0.000051600968],"about_ca_topic_score_codex":0.0019572936,"about_ca_topic_score_gemma":0.0018064386,"teacher_disagreement_score":0.008051501,"about_ca_system_score_codex":0.0013896452,"about_ca_system_score_gemma":0.00090232794,"threshold_uncertainty_score":0.042580903},"labels":[],"label_agreement":null},{"id":"W2945245827","doi":"10.1109/spin.2019.8711636","title":"Accuracy of Convolution Neural Networks for Classifying Sentiments on Movie Reviews","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"","keywords":"Computer science; Sentence; Word embedding; Convolutional neural network; Artificial intelligence; Kernel (algebra); Embedding; Word (group theory); Convolution (computer science); Natural language processing; Task (project management); Artificial neural network; Pattern recognition (psychology); Machine learning; Mathematics","score_opus":0.06830139006968847,"score_gpt":0.31361457101722306,"score_spread":0.2453131809475346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945245827","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95086443,0.0026093465,0.029407606,0.0007559941,0.00038215274,0.00005612138,0.0025989432,0.0018356893,0.011489729],"genre_scores_gemma":[0.98518884,0.00044403833,0.007962365,0.00006655143,0.000065987944,0.000014600683,0.0037007274,0.00006607604,0.002490896],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986688,0.000333935,0.00015608424,0.00030126737,0.00036947342,0.00017057943],"domain_scores_gemma":[0.99592197,0.0020719373,0.00037786888,0.00045221858,0.0010109547,0.00016498976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032873205,0.0013403223,0.0005847377,0.0015809617,0.00031135618,0.0012053554,0.0005157148,0.0009780243,0.0018265245],"category_scores_gemma":[0.010769068,0.00023452009,0.0006209082,0.0008177594,0.0002920768,0.0016534957,0.00063572114,0.0006965921,0.0014765806],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025104315,0.00047238916,0.19388795,0.0005849589,0.0009826258,0.00025993647,0.00034623806,0.13646606,0.0274778,0.0018158912,0.028976662,0.6062192],"study_design_scores_gemma":[0.00002452028,0.00018542985,0.039805055,0.000051913783,0.00010619265,0.000092219016,0.00015894315,0.9444504,0.01250542,0.0010412327,0.0015463447,0.000032369975],"about_ca_topic_score_codex":0.009818834,"about_ca_topic_score_gemma":0.011274112,"teacher_disagreement_score":0.009818834,"about_ca_system_score_codex":0.0008099233,"about_ca_system_score_gemma":0.0004753769,"threshold_uncertainty_score":0.019523323},"labels":[],"label_agreement":null},{"id":"W2945699119","doi":"10.1109/infoteh.2019.8717655","title":"Automatic Text Summarization of News Articles in Serbian Language","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Automatic summarization; Computer science; Text graph; Multi-document summarization; Information retrieval; Natural language processing; Reading (process); Graph; Process (computing); Serbian; Artificial intelligence; World Wide Web; Linguistics","score_opus":0.010495716490630273,"score_gpt":0.2305994945684212,"score_spread":0.2201037780777909,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945699119","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.83327526,0.012173751,0.090546794,0.001071731,0.00086043385,0.001372396,0.031874385,0.015870964,0.012954307],"genre_scores_gemma":[0.621347,0.0036112324,0.23026688,0.00020128702,0.00071671273,0.0008383784,0.12894861,0.00063114887,0.013438826],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99908423,0.00032484342,0.0001105766,0.00018752058,0.00021617992,0.00007669059],"domain_scores_gemma":[0.9965419,0.0012816539,0.00046178227,0.000261173,0.0013195085,0.00013402416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012901268,0.0011629176,0.0009729563,0.006670179,0.0007813471,0.0014949249,0.0005509444,0.00044349715,0.0018703431],"category_scores_gemma":[0.004249824,0.00023543372,0.00057238754,0.0028233016,0.00018976971,0.0009569283,0.00050700695,0.00055819756,0.002072265],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016768279,0.00077210733,0.014962608,0.0035648877,0.0003343295,0.0012306834,0.0021981588,0.010413911,0.101388335,0.0010421273,0.054722305,0.8076938],"study_design_scores_gemma":[0.000637905,0.0037947414,0.22287467,0.0008283652,0.001205434,0.0021530774,0.01037526,0.3219312,0.17836699,0.004469337,0.25307515,0.0002878811],"about_ca_topic_score_codex":0.0034170584,"about_ca_topic_score_gemma":0.008126434,"teacher_disagreement_score":0.006670179,"about_ca_system_score_codex":0.00036743138,"about_ca_system_score_gemma":0.0006550545,"threshold_uncertainty_score":0.0068229437},"labels":[],"label_agreement":null},{"id":"W2945702027","doi":"10.1007/978-3-030-18305-9_12","title":"Efficient Transformer-Based Sentence Encoding for Sentence Pair Modelling","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Sentence; Computer science; Encoder; Natural language processing; Transformer; Paraphrase; Artificial intelligence; Semantic similarity; Textual entailment; Atomic sentence; Speech recognition; Logical consequence","score_opus":0.0350367527427273,"score_gpt":0.24631197798119744,"score_spread":0.21127522523847014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945702027","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053058867,0.00030233705,0.9786147,0.00019080074,0.00018517661,0.00018454458,0.0023264845,0.010313459,0.0025766762],"genre_scores_gemma":[0.2157103,0.00062989944,0.76143056,0.00021812592,0.0002395231,0.00041330935,0.013244442,0.0023762581,0.005737529],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987728,0.00037792537,0.00014667977,0.00029195411,0.00028990326,0.00012065721],"domain_scores_gemma":[0.9975733,0.0012160481,0.0001031291,0.00044152886,0.00056862313,0.00009728978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011532442,0.0012576843,0.0012036837,0.001404831,0.0006701589,0.0020472442,0.0024178468,0.0011143438,0.01834035],"category_scores_gemma":[0.0049209064,0.00067595596,0.0017627011,0.0017876313,0.00047200872,0.004525709,0.0024150065,0.0022359341,0.009994326],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010703428,0.0002649443,0.0005834081,0.0010259878,0.00015098852,0.00054275064,0.0005339848,0.03992567,0.041078813,0.07685701,0.051528003,0.7864381],"study_design_scores_gemma":[0.00007315609,0.00014947045,0.00023678513,0.00008121709,0.00013106481,0.00040907593,0.00017941483,0.8426145,0.034125216,0.10096198,0.020974554,0.00006353905],"about_ca_topic_score_codex":0.0025137868,"about_ca_topic_score_gemma":0.0037853802,"teacher_disagreement_score":0.01834035,"about_ca_system_score_codex":0.0009855293,"about_ca_system_score_gemma":0.0015928517,"threshold_uncertainty_score":0.061354578},"labels":[],"label_agreement":null},{"id":"W2946008212","doi":"","title":"Copy this Sentence.","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Sequence (biology); Copying; Computer science; Sentence; Task (project management); Set (abstract data type); Sequence learning; Artificial intelligence; Implementation; Simple (philosophy); Natural language processing; Theoretical computer science; Programming language","score_opus":0.08317776917945031,"score_gpt":0.18466362078067056,"score_spread":0.10148585160122024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2946008212","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0112158,0.0026572216,0.02978754,0.02597047,0.029741704,0.00066732056,0.121144295,0.027685203,0.7511304],"genre_scores_gemma":[0.06286226,0.001609569,0.014108697,0.008221113,0.0038600436,0.00038677297,0.06780796,0.009624567,0.83151907],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999728,0.00004405908,0.000024236551,0.000063564745,0.00010858082,0.000031539657],"domain_scores_gemma":[0.9986558,0.00037337342,0.00007337156,0.00029941145,0.00045193871,0.00014614403],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00044485126,0.0009688212,0.00057539955,0.0009305899,0.00083983596,0.0021284767,0.00086888135,0.0012771708,0.60157883],"category_scores_gemma":[0.005853539,0.0003047597,0.00055345555,0.00076311664,0.00033796765,0.0031574129,0.0021799318,0.0013001517,0.44549167],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009669105,0.00002593166,0.00025503067,0.00018477859,0.000013202714,0.0001925021,0.00018520764,0.00008883292,0.0010051931,0.00501054,0.94722474,0.0457173],"study_design_scores_gemma":[0.000019427502,0.000028790624,0.0007956085,0.00006405755,0.000009963921,0.00026713262,0.00015444269,0.0005888761,0.0011993846,0.0073261624,0.98952556,0.0000205101],"about_ca_topic_score_codex":0.0018315817,"about_ca_topic_score_gemma":0.003046511,"teacher_disagreement_score":0.39842117,"about_ca_system_score_codex":0.00061385735,"about_ca_system_score_gemma":0.00050645147,"threshold_uncertainty_score":0.5682994},"labels":[],"label_agreement":null},{"id":"W2946869570","doi":"10.48550/arxiv.1905.11912","title":"A Cross-Domain Transferable Neural Coherence Model","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Discriminative model; Coherence (philosophical gambling strategy); Computer science; Generalization; Artificial intelligence; Benchmark (surveying); Sentence; Readability; Domain (mathematical analysis); Generative grammar; Space (punctuation); Generative model; Machine learning; Natural language processing; Mathematics","score_opus":0.0924450073983217,"score_gpt":0.20407974256581465,"score_spread":0.11163473516749295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2946869570","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.115305975,0.0008938068,0.87147987,0.0012911543,0.000099226905,0.00012429272,0.00089546334,0.0017587744,0.008151487],"genre_scores_gemma":[0.9032283,0.00036630032,0.07908382,0.0005500696,0.0001317786,0.00028122286,0.001742639,0.0002173524,0.014398459],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995023,0.00012757191,0.00002466243,0.00021283256,0.00007674009,0.00005590254],"domain_scores_gemma":[0.99891245,0.00049190345,0.00013149074,0.0001644663,0.00023868782,0.000061037485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010374142,0.0005323379,0.0005247264,0.0008046401,0.00038934525,0.00072831486,0.0017759943,0.0012607217,0.0032577717],"category_scores_gemma":[0.004261048,0.00037901872,0.0006362309,0.0009826157,0.00076663063,0.0027120174,0.0014753971,0.0017328383,0.0008499269],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041910625,0.00039781863,0.0060534896,0.00024957632,0.000201996,0.00034974984,0.00084789534,0.54438746,0.01812483,0.08200318,0.014655772,0.3323092],"study_design_scores_gemma":[0.000019243436,0.000051247996,0.0007101827,0.000009407176,0.00002315248,0.000048767895,0.000026920941,0.9795944,0.00086063,0.017244617,0.0014019072,0.000009493317],"about_ca_topic_score_codex":0.0053133294,"about_ca_topic_score_gemma":0.008529317,"teacher_disagreement_score":0.0053133294,"about_ca_system_score_codex":0.000895097,"about_ca_system_score_gemma":0.0007103954,"threshold_uncertainty_score":0.010898352},"labels":[],"label_agreement":null},{"id":"W2947119606","doi":"","title":"H2oloo at TREC 2018: Cross-Collection Relevance Transfer for the Common Core Track.","year":2018,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Track (disk drive); Relevance (law); Core (optical fiber); Transfer (computing); Information retrieval; Telecommunications; Parallel computing","score_opus":0.0819327481816418,"score_gpt":0.3220630675340519,"score_spread":0.24013031935241008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2947119606","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15944344,0.018958833,0.16079547,0.008221842,0.019098647,0.009272477,0.41740757,0.14498474,0.06181698],"genre_scores_gemma":[0.1284813,0.0011410926,0.121823244,0.0015776983,0.0017248151,0.0032318395,0.68185186,0.0053853206,0.05478282],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99356365,0.002311767,0.0003134765,0.0013181904,0.0016082273,0.00088463764],"domain_scores_gemma":[0.9891642,0.0022777407,0.00024757642,0.002822126,0.004147489,0.0013407879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011126602,0.0032375613,0.0022974887,0.004455529,0.0037785694,0.0032995248,0.0038958848,0.0030023055,0.021720802],"category_scores_gemma":[0.022782473,0.0009350904,0.0016886436,0.0034944874,0.0010516014,0.005731988,0.0059572835,0.004136692,0.016870053],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014773731,0.0007355039,0.002120504,0.0008959327,0.00032425422,0.0001077829,0.00030691785,0.0030472507,0.010279685,0.00089205184,0.86599815,0.11381455],"study_design_scores_gemma":[0.0061199656,0.0032509,0.041057084,0.0005776383,0.0013599327,0.0006725965,0.002023429,0.16256784,0.0746528,0.013746621,0.69312406,0.00084706995],"about_ca_topic_score_codex":0.07219913,"about_ca_topic_score_gemma":0.15303023,"teacher_disagreement_score":0.07219913,"about_ca_system_score_codex":0.0027739517,"about_ca_system_score_gemma":0.008280939,"threshold_uncertainty_score":0.14355779},"labels":[],"label_agreement":null},{"id":"W2947143940","doi":"10.1177/1747021819855621","title":"On the role of autobiographical knowledge in shaping belief in the future occurrence of imagined events","year":2019,"lang":"en","type":"article","venue":"Quarterly Journal of Experimental Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"Social Sciences and Humanities Research Council of Canada; Fonds De La Recherche Scientifique - FNRS; European Commission","keywords":"Psychology; Autobiographical memory; Cognitive psychology; Cognitive science; Communication; Recall","score_opus":0.01875171288334468,"score_gpt":0.32349316902456643,"score_spread":0.3047414561412217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2947143940","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9904004,0.00014607685,0.004458666,0.00028931524,0.0000051576258,0.0000149899,0.000017531805,0.000008718384,0.0046590604],"genre_scores_gemma":[0.99896824,0.00007622278,0.0007792654,0.000019232075,0.0000018600933,0.0000052366395,0.000019400464,0.0000029951198,0.00012744576],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.999271,0.00041082056,0.000031312982,0.00009835524,0.00012537895,0.00006314619],"domain_scores_gemma":[0.98568535,0.010339803,0.0016046865,0.0012355854,0.00057823304,0.0005562627],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028144843,0.00019171754,0.00018006114,0.00061468076,0.00037674955,0.0020041037,0.00033359,0.00060378877,0.0021165048],"category_scores_gemma":[0.015660105,0.00026103196,0.00025227113,0.00033120473,0.0016802978,0.0018617547,0.0009554546,0.00096362404,0.0001323657],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021444247,0.0014790732,0.5833414,0.00064300816,0.0004361758,0.001346804,0.10153264,0.0076317866,0.06095029,0.06364771,0.0007632008,0.17608342],"study_design_scores_gemma":[0.00010313139,0.0008542529,0.8502854,0.00039328155,0.0003505516,0.001002492,0.030542918,0.03134457,0.014463199,0.06487681,0.0056351423,0.00014827611],"about_ca_topic_score_codex":0.0014168087,"about_ca_topic_score_gemma":0.002237491,"teacher_disagreement_score":0.0028144843,"about_ca_system_score_codex":0.0005371824,"about_ca_system_score_gemma":0.00052097347,"threshold_uncertainty_score":0.014884591},"labels":[],"label_agreement":null},{"id":"W2948710720","doi":"10.18653/v1/p19-1004","title":"Do Neural Dialog Systems Use the Conversation History Effectively? An Empirical Study","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"Nvidia","keywords":"Dialog box; Computer science; Generative grammar; Flexibility (engineering); Shuffling; Context (archaeology); Artificial intelligence; Transformer; Code (set theory); Conversation; Dialog system; Natural language processing; Machine learning; Human–computer interaction; Programming language; Communication; World Wide Web; Psychology; Engineering","score_opus":0.10604917870750143,"score_gpt":0.30926652051693976,"score_spread":0.20321734180943835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2948710720","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9753044,0.0023978804,0.015716042,0.0007664411,0.00006127753,0.00012825635,0.001249425,0.00029582335,0.0040804883],"genre_scores_gemma":[0.9941855,0.00031233273,0.003395483,0.00008804819,0.000034003926,0.000049507602,0.0013467987,0.000049145998,0.0005391579],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9894404,0.00742089,0.0005098988,0.0015366825,0.0007961684,0.0002959506],"domain_scores_gemma":[0.8796498,0.099583544,0.0051137605,0.01099752,0.0035338139,0.0011215201],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01623791,0.000837775,0.0006917346,0.0011284956,0.0008965142,0.0017950892,0.0013617981,0.0014221321,0.0022281087],"category_scores_gemma":[0.09303408,0.00080657884,0.00063911773,0.0010222591,0.0016558085,0.005492673,0.0016345428,0.0026852502,0.0011636484],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0058422475,0.002387303,0.37853056,0.0027415913,0.002145136,0.0006613047,0.008887505,0.21614726,0.018964775,0.007882627,0.014446675,0.34136307],"study_design_scores_gemma":[0.00025259893,0.00156905,0.22397864,0.00023932515,0.0004995986,0.0010622912,0.002946074,0.7368421,0.010809635,0.013836398,0.007729324,0.00023499831],"about_ca_topic_score_codex":0.0046852496,"about_ca_topic_score_gemma":0.0053569963,"teacher_disagreement_score":0.01623791,"about_ca_system_score_codex":0.0012872313,"about_ca_system_score_gemma":0.0004678395,"threshold_uncertainty_score":0.08587533},"labels":[],"label_agreement":null},{"id":"W2949034987","doi":"10.18653/v1/p19-1153","title":"Towards Lossless Encoding of Sentences","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Compute Canada","keywords":"Lossless compression; Sentence; Computer science; Encoding (memory); Embedding; Feature (linguistics); Natural language processing; Focus (optics); Task (project management); Artificial intelligence; Sequence labeling; Compression (physics); Sequence (biology); Field (mathematics); Speech recognition; Data compression; Linguistics; Mathematics","score_opus":0.05661246926732516,"score_gpt":0.28877376260446097,"score_spread":0.2321612933371358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949034987","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018804284,0.0006285949,0.97504914,0.00061351317,0.00020902499,0.000046911136,0.0007794959,0.0020347442,0.0018343993],"genre_scores_gemma":[0.34473497,0.0016775472,0.63029253,0.0012704269,0.0006541968,0.0004817166,0.0054939287,0.0012006471,0.014194048],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99899334,0.0003778544,0.00008120037,0.00019654587,0.00027167416,0.000079452766],"domain_scores_gemma":[0.9980242,0.0007672873,0.00016171188,0.0006085351,0.00036658414,0.0000717584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00087919063,0.0008459952,0.00061562937,0.00089021743,0.00025368916,0.0010260381,0.0007866083,0.0007557624,0.0032315734],"category_scores_gemma":[0.006543089,0.00030035625,0.0004917859,0.00078487274,0.0006480952,0.003371688,0.0015757266,0.0017693073,0.0023897032],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075787853,0.00021632147,0.0007919,0.0004391605,0.00007325368,0.00036830208,0.0005945611,0.089405626,0.05965897,0.09561447,0.033603612,0.718476],"study_design_scores_gemma":[0.00005992281,0.00028935273,0.00043960602,0.00009232789,0.000043826,0.0003653482,0.0001679055,0.7894003,0.04070337,0.14134239,0.027054515,0.000041194126],"about_ca_topic_score_codex":0.0005305022,"about_ca_topic_score_gemma":0.0007980874,"teacher_disagreement_score":0.0032315734,"about_ca_system_score_codex":0.00045143263,"about_ca_system_score_gemma":0.0005814673,"threshold_uncertainty_score":0.010810733},"labels":[],"label_agreement":null},{"id":"W2949223751","doi":"10.48550/arxiv.1702.03470","title":"Vector Embedding of Wikipedia Concepts and Entities","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Embedding; Word embedding; Natural language processing; Similarity (geometry); Popularity; Artificial intelligence; Analogy; Word (group theory); Task (project management); Deep learning; Information retrieval; Linguistics; Image (mathematics)","score_opus":0.08734957569459353,"score_gpt":0.22301074969636275,"score_spread":0.13566117400176922,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949223751","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15845644,0.0037257925,0.817985,0.00058012217,0.00056975463,0.00026352244,0.0066610556,0.003821274,0.007936956],"genre_scores_gemma":[0.6827926,0.0017872095,0.28970748,0.0001598405,0.00020613061,0.0002731825,0.017190043,0.00020407427,0.007679588],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99943894,0.00015205969,0.00004572719,0.00019886016,0.00011856309,0.00004582629],"domain_scores_gemma":[0.9990927,0.0003172573,0.00011667549,0.00017696328,0.00025557308,0.00004074991],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00054331456,0.00092794665,0.000490725,0.002553425,0.00023166543,0.0008419485,0.0006632724,0.00062265916,0.0020167418],"category_scores_gemma":[0.0030448502,0.00023754446,0.0006485271,0.0023352746,0.00030414783,0.0023521024,0.00081204355,0.00078098156,0.0009506381],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041065263,0.00038324555,0.0075559504,0.0011376695,0.00038335577,0.00035067144,0.00046773025,0.10828734,0.024979973,0.029637307,0.03182507,0.794581],"study_design_scores_gemma":[0.000027794089,0.00014975773,0.004265823,0.000079806225,0.00008318732,0.00028570634,0.00018615466,0.9354877,0.013978233,0.025038406,0.020369358,0.00004797728],"about_ca_topic_score_codex":0.0034407745,"about_ca_topic_score_gemma":0.005048891,"teacher_disagreement_score":0.0034407745,"about_ca_system_score_codex":0.00045741818,"about_ca_system_score_gemma":0.0005145062,"threshold_uncertainty_score":0.0068414807},"labels":[],"label_agreement":null},{"id":"W2949400899","doi":"10.48550/arxiv.1308.6300","title":"Computing Lexical Contrast","year":2013,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; National Research Council Canada","funders":"","keywords":"Contrast (vision); Meaning (existential); Word (group theory); Computer science; Natural language processing; Artificial intelligence; Lexical item; Mathematics; Linguistics; Psychology; Philosophy","score_opus":0.08929011440420782,"score_gpt":0.18658307879358393,"score_spread":0.09729296438937611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949400899","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57115465,0.0028721537,0.3219409,0.0015351335,0.0009122012,0.00059616414,0.012479012,0.0048328643,0.083676964],"genre_scores_gemma":[0.8708991,0.00036190273,0.113939784,0.00027999326,0.00021513758,0.00029465923,0.010497791,0.00034780684,0.0031639074],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957029,0.00062999706,0.0005369831,0.0016403944,0.0010625084,0.00042718914],"domain_scores_gemma":[0.99119484,0.0046408335,0.0006331609,0.0011779089,0.0018631333,0.00048999675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019625654,0.0010049433,0.0014951534,0.011022285,0.0019477829,0.0054299016,0.001224971,0.001557829,0.012621056],"category_scores_gemma":[0.021805493,0.00064052996,0.0010915124,0.0064030746,0.0011888037,0.010278643,0.003831242,0.0014770074,0.003939301],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025004165,0.00055836886,0.11782805,0.0014582372,0.0005718485,0.0021240185,0.0023435436,0.010234084,0.04334768,0.13431546,0.02971517,0.655003],"study_design_scores_gemma":[0.00037998284,0.00089563173,0.07312678,0.00038781666,0.0007083507,0.0045967093,0.0055051334,0.2148697,0.038380284,0.5647683,0.09605747,0.0003238155],"about_ca_topic_score_codex":0.0016879025,"about_ca_topic_score_gemma":0.0018364526,"teacher_disagreement_score":0.012621056,"about_ca_system_score_codex":0.0013345134,"about_ca_system_score_gemma":0.0011182731,"threshold_uncertainty_score":0.042221606},"labels":[],"label_agreement":null},{"id":"W2949423690","doi":"10.18653/v1/p19-1067","title":"A Cross-Domain Transferable Neural Coherence Model","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Canadian Institute for Advanced Research; McGill University","funders":"","keywords":"Coherence (philosophical gambling strategy); Computer science; Domain (mathematical analysis); Cognitive science; Natural language processing; Artificial intelligence; Physics; Psychology; Mathematics; Quantum mechanics; Mathematical analysis","score_opus":0.046871082728640094,"score_gpt":0.28972156300119944,"score_spread":0.24285048027255934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949423690","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037322946,0.00171656,0.9506159,0.0010475044,0.0001703243,0.00008188038,0.00069606694,0.0018795816,0.006469247],"genre_scores_gemma":[0.82850754,0.00088498264,0.1483956,0.000477426,0.0002090503,0.00028662925,0.0021119574,0.00044836634,0.01867849],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99964786,0.00011454242,0.000017426259,0.00013077812,0.000045975874,0.00004343209],"domain_scores_gemma":[0.99922204,0.00039284062,0.00006257111,0.000097885655,0.00018186313,0.000042809752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010979741,0.0005722157,0.0006866332,0.0007376081,0.0004061685,0.0009740508,0.0017490059,0.001142739,0.00522125],"category_scores_gemma":[0.0030857273,0.0005167986,0.00078877044,0.0011491199,0.00044111835,0.0026941628,0.0015659768,0.0017196146,0.0013303244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006302039,0.00033319683,0.002207499,0.000271993,0.0004348946,0.00019697667,0.0003536658,0.43617794,0.01067361,0.05999342,0.015921196,0.4728054],"study_design_scores_gemma":[0.000018847128,0.000036325186,0.0002905933,0.000008868417,0.000031965126,0.000019630084,0.000013479021,0.9868601,0.0004389806,0.011286191,0.0009868373,0.0000081086],"about_ca_topic_score_codex":0.008232278,"about_ca_topic_score_gemma":0.010477344,"teacher_disagreement_score":0.008232278,"about_ca_system_score_codex":0.00071778113,"about_ca_system_score_gemma":0.0007886705,"threshold_uncertainty_score":0.017466843},"labels":[],"label_agreement":null},{"id":"W2949446888","doi":"10.48550/arxiv.1807.10805","title":"Improving Neural Sequence Labelling using Additional Linguistic Information","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Chunking (psychology); Computer science; Natural language processing; Sequence labeling; Artificial intelligence; Sequence (biology); Named-entity recognition; Word (group theory); Labelling; Sentence; Benchmark (surveying); Task (project management); Linguistics","score_opus":0.11366319767681857,"score_gpt":0.20310247414074956,"score_spread":0.08943927646393099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949446888","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041421853,0.0010784732,0.94639623,0.0006129439,0.00024864922,0.00010557408,0.0005623175,0.0057457197,0.003828206],"genre_scores_gemma":[0.5161473,0.0009659034,0.45718402,0.0010157737,0.0003017704,0.00042415163,0.0065799328,0.0008103182,0.016570855],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992638,0.00020623438,0.000048163838,0.0002951707,0.000120643796,0.00006593527],"domain_scores_gemma":[0.99650335,0.0018438783,0.00021229497,0.00064134685,0.0006776816,0.000121385216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015055736,0.0012630954,0.0011688665,0.00131361,0.00067874754,0.0010922124,0.0021750613,0.0019904645,0.0033594223],"category_scores_gemma":[0.0071252594,0.0005764672,0.0010788486,0.0013752031,0.00067932054,0.004755227,0.0014931465,0.002605011,0.0032675285],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035767598,0.0004240748,0.00281117,0.00036666347,0.00011330398,0.00020995247,0.0004197822,0.28835252,0.018143324,0.014017982,0.017902985,0.6568805],"study_design_scores_gemma":[0.000011751548,0.00004021535,0.0002482003,0.000022199898,0.000021926055,0.00004403707,0.000030015468,0.9808882,0.0032129798,0.013189691,0.002277097,0.000013741193],"about_ca_topic_score_codex":0.0070205824,"about_ca_topic_score_gemma":0.014260165,"teacher_disagreement_score":0.0070205824,"about_ca_system_score_codex":0.0012348255,"about_ca_system_score_gemma":0.0016639711,"threshold_uncertainty_score":0.013959467},"labels":[],"label_agreement":null},{"id":"W2949568611","doi":"10.18653/v1/p19-1134","title":"Fine-tuning Pre-Trained Transformer Language Models to Distantly Supervised Relation Extraction","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung; Bundesministerium für Verkehr und Digitale Infrastruktur","keywords":"Relationship extraction; Transformer; Computer science; Artificial intelligence; Relation (database); Generative grammar; Natural language processing; Machine learning; Recall; Set (abstract data type); Language model; F1 score; Training set; Generative model; Information extraction; Data mining; Engineering; Linguistics","score_opus":0.03763948805944712,"score_gpt":0.2848659437693302,"score_spread":0.2472264557098831,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949568611","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1639929,0.002119578,0.79333067,0.0009516212,0.00030490282,0.00039913596,0.003474185,0.028246507,0.0071805194],"genre_scores_gemma":[0.68344504,0.0005318014,0.2865558,0.0008144231,0.00018597954,0.00040763235,0.019075956,0.0010957479,0.007887576],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99829096,0.00046596088,0.00010952343,0.00077629345,0.00021012651,0.00014717413],"domain_scores_gemma":[0.9958462,0.0027671969,0.00016375144,0.0006160678,0.00047177213,0.00013511318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002569637,0.0021502124,0.0012077827,0.00228551,0.00068260176,0.0016468139,0.0031445157,0.0014468196,0.00311008],"category_scores_gemma":[0.006840651,0.0007218578,0.0024210012,0.0017518764,0.00086642464,0.0043896013,0.0020873107,0.0033927325,0.0037528619],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082105293,0.0008786613,0.011142022,0.0005460771,0.00043971147,0.0006198448,0.0006598044,0.2640729,0.018915975,0.009756868,0.029962528,0.6621845],"study_design_scores_gemma":[0.000032375174,0.000062957886,0.00045799065,0.000016475708,0.00004226176,0.000101742546,0.000053156702,0.988538,0.004058166,0.004938498,0.0016834842,0.000014820695],"about_ca_topic_score_codex":0.009307788,"about_ca_topic_score_gemma":0.020343892,"teacher_disagreement_score":0.009307788,"about_ca_system_score_codex":0.0013174019,"about_ca_system_score_gemma":0.0017239683,"threshold_uncertainty_score":0.018507242},"labels":[],"label_agreement":null},{"id":"W2950021574","doi":"10.1093/bioinformatics/btz504","title":"Towards reliable named entity recognition in the biomedical domain","year":2019,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Topic Modeling","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Human Genome Research Institute; Compute Canada; National Institutes of Health; Nvidia","keywords":"CRFS; Conditional random field; Computer science; Artificial intelligence; Dropout (neural networks); Transfer of learning; Named-entity recognition; Natural language processing; Machine learning; Task (project management); Regularization (linguistics); Deep learning; Sequence labeling; Source code; Generalization; Multi-task learning","score_opus":0.021713303600085557,"score_gpt":0.2395539452077731,"score_spread":0.21784064160768754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950021574","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036250085,0.0063629868,0.876975,0.004993779,0.0007837351,0.00030144415,0.015919266,0.05203746,0.0063761743],"genre_scores_gemma":[0.22057045,0.0025695087,0.7038203,0.0014044945,0.0005146472,0.0003308952,0.063585676,0.0018262522,0.005377881],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939559,0.0023842633,0.0004464016,0.0016652675,0.0012942689,0.00025403078],"domain_scores_gemma":[0.97684413,0.011695658,0.0016138775,0.0050251214,0.0041815387,0.0006396629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012952396,0.002213101,0.001525546,0.004336638,0.000815239,0.0029950151,0.0030900503,0.003548447,0.0060580876],"category_scores_gemma":[0.035673484,0.00079220725,0.0014603928,0.0038502007,0.0012622795,0.009502488,0.0047160834,0.0037406995,0.011964508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011951517,0.0004039822,0.00914653,0.0029715921,0.00038750563,0.00076299737,0.00062363,0.11617828,0.034305066,0.02638712,0.13595565,0.6716825],"study_design_scores_gemma":[0.00008797746,0.00018678725,0.0035900967,0.00038959555,0.00013646789,0.0006065902,0.00038211732,0.8132056,0.06442299,0.049370233,0.06749763,0.00012385815],"about_ca_topic_score_codex":0.0031115648,"about_ca_topic_score_gemma":0.0026875762,"teacher_disagreement_score":0.012952396,"about_ca_system_score_codex":0.00095546566,"about_ca_system_score_gemma":0.0023344816,"threshold_uncertainty_score":0.068499684},"labels":[],"label_agreement":null},{"id":"W2950072837","doi":"10.48550/arxiv.1711.07632","title":"Generating Thematic Chinese Poetry using Conditional Variational Autoencoders with Hybrid Decoders","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Autoencoder; Computer science; Word2vec; Poetry; Context (archaeology); Artificial intelligence; Relevance (law); Natural language processing; Theme (computing); Sequence (biology); Artificial neural network; Theoretical computer science; Pattern recognition (psychology); Linguistics; History; Philosophy","score_opus":0.07932672630128064,"score_gpt":0.2133251980187364,"score_spread":0.13399847171745577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950072837","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04758493,0.00022020665,0.9483906,0.00022449825,0.00006174696,0.000043904744,0.00012527089,0.0009661642,0.0023826042],"genre_scores_gemma":[0.7215012,0.00025622023,0.268956,0.00020888477,0.000071390496,0.00014650922,0.0007718675,0.00025566044,0.007832309],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997454,0.000092032446,0.000014006511,0.00007182527,0.000047756825,0.000028828379],"domain_scores_gemma":[0.9993629,0.00039944478,0.000032834203,0.00006287334,0.000114200804,0.000027644635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00064192736,0.0005758442,0.00046215398,0.00032117128,0.00020738124,0.00048794923,0.0007109261,0.0005811905,0.0019373788],"category_scores_gemma":[0.0020687506,0.00036876413,0.0006764257,0.00032058728,0.00041315478,0.00080051715,0.0007227544,0.00096766517,0.000693503],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016381234,0.00013530465,0.0014599084,0.0001448465,0.00010773506,0.00015944031,0.00021005115,0.6359129,0.017693361,0.020725869,0.0037324857,0.31955433],"study_design_scores_gemma":[0.0000048399,0.000010283219,0.000060395367,0.0000025674015,0.0000041626126,0.000010983824,0.000006624602,0.99608105,0.0016678232,0.0018649881,0.00028384093,0.0000025374975],"about_ca_topic_score_codex":0.0029001434,"about_ca_topic_score_gemma":0.0039886944,"teacher_disagreement_score":0.0029001434,"about_ca_system_score_codex":0.0003564303,"about_ca_system_score_gemma":0.0005658077,"threshold_uncertainty_score":0.0064812303},"labels":[],"label_agreement":null},{"id":"W2950632837","doi":"","title":"Measuring Semantic Similarity by Latent Relational Analysis","year":2005,"lang":"en","type":"preprint","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Similarity (geometry); Analogy; Natural language processing; Semantic similarity; Latent semantic analysis; Computer science; Artificial intelligence; Word (group theory); Noun; Cosine similarity; Vector space; Vector space model; Statistical relational learning; Relational database; Mathematics; Pattern recognition (psychology); Data mining; Linguistics; Pure mathematics","score_opus":0.06573422506007907,"score_gpt":0.25327344331676105,"score_spread":0.18753921825668196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950632837","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041014105,0.001230673,0.94987893,0.0002953356,0.000054044987,0.00023365018,0.001425391,0.0012063335,0.004661547],"genre_scores_gemma":[0.53571177,0.0009069093,0.4567427,0.00012878979,0.00018262731,0.000567901,0.004291634,0.00025914604,0.0012085512],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9888433,0.0041695107,0.00093253097,0.0020190685,0.0037631604,0.00027237143],"domain_scores_gemma":[0.984364,0.0090211155,0.0021946256,0.002462674,0.0016427262,0.00031495886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053934315,0.0009925602,0.0013074615,0.012261257,0.0010199609,0.0043757404,0.0015923202,0.0012320998,0.003918427],"category_scores_gemma":[0.031096265,0.0005023228,0.00174901,0.009175377,0.0014478031,0.009106883,0.0032271529,0.0015178741,0.0014933662],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064964127,0.00061258825,0.04150473,0.001204472,0.001027118,0.00026103866,0.0030875094,0.03373081,0.019545013,0.17530203,0.007937388,0.71513766],"study_design_scores_gemma":[0.00009925585,0.00035407883,0.028647399,0.0001918701,0.00039723038,0.00070958206,0.00171811,0.528301,0.011512151,0.40837502,0.019440006,0.00025428113],"about_ca_topic_score_codex":0.0019595677,"about_ca_topic_score_gemma":0.0016242957,"teacher_disagreement_score":0.012261257,"about_ca_system_score_codex":0.0013010603,"about_ca_system_score_gemma":0.00108672,"threshold_uncertainty_score":0.028523505},"labels":[],"label_agreement":null},{"id":"W2950695840","doi":"10.18653/v1/p19-1486","title":"Simple and Effective Curriculum Pointer-Generator Networks for Reading Comprehension over Long Narratives","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Pointer (user interface); Computer science; Narrative; Generator (circuit theory); Simple (philosophy); Zhàng; Reading comprehension; Linguistics; Reading (process); Mathematics education; Artificial intelligence; Mathematics; History; Philosophy; Physics; China; Epistemology","score_opus":0.015157624067037783,"score_gpt":0.27467499757980585,"score_spread":0.25951737351276805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950695840","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12668623,0.000793311,0.8481537,0.0012157323,0.000112480884,0.0003173969,0.0011446852,0.012474579,0.009102033],"genre_scores_gemma":[0.78724456,0.00030411314,0.20469415,0.00011080811,0.00006022953,0.000370115,0.0017130502,0.00057447836,0.0049284613],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994411,0.0002134586,0.00003293814,0.0002131382,0.000051518884,0.000047953472],"domain_scores_gemma":[0.9964354,0.0023331002,0.00018041769,0.0005094443,0.0003477821,0.00019386048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014450523,0.0008646642,0.0005113698,0.0010368619,0.000638543,0.0012526917,0.001592089,0.001154169,0.010496013],"category_scores_gemma":[0.012171993,0.00053004606,0.000515658,0.00064850703,0.00052583264,0.005664585,0.0021059667,0.0014772852,0.0022060175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006182223,0.00041565511,0.0059679355,0.00040522052,0.000097903954,0.00019019835,0.001239289,0.061561804,0.016196802,0.042929016,0.014489338,0.8558887],"study_design_scores_gemma":[0.00007825913,0.000114084716,0.001463401,0.000054506072,0.000077544726,0.00006823422,0.00024940315,0.8950457,0.008330045,0.09045579,0.004036861,0.000026157984],"about_ca_topic_score_codex":0.0035998418,"about_ca_topic_score_gemma":0.005964706,"teacher_disagreement_score":0.010496013,"about_ca_system_score_codex":0.0010851745,"about_ca_system_score_gemma":0.0013720419,"threshold_uncertainty_score":0.03511268},"labels":[],"label_agreement":null},{"id":"W2950824039","doi":"","title":"The Latent Relation Mapping Engine: Algorithm and Experiments","year":2008,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":107,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Analogy; Computer science; Relation (database); Set (abstract data type); Latent semantic analysis; Variety (cybernetics); Core (optical fiber); Natural language processing; Artificial intelligence; Theoretical computer science; Algorithm; Data mining; Linguistics; Programming language","score_opus":0.032613120169892845,"score_gpt":0.2335917505973968,"score_spread":0.20097863042750397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950824039","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57817584,0.0045625027,0.35438067,0.002219222,0.00056352763,0.002873377,0.0037895062,0.02733284,0.026102435],"genre_scores_gemma":[0.51813895,0.0007618712,0.470135,0.0004229549,0.000075428165,0.0017057117,0.0041709277,0.0009301207,0.0036590488],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972071,0.0011998381,0.00024136678,0.0005853459,0.0005611234,0.00020530788],"domain_scores_gemma":[0.97832936,0.01745503,0.00029973633,0.0019580128,0.0016536195,0.0003041724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006044439,0.0015015261,0.0015360058,0.0013540405,0.00086805795,0.0014857484,0.0031416514,0.0030587616,0.008462623],"category_scores_gemma":[0.029244414,0.0006094255,0.0007098185,0.0024684726,0.00069490063,0.004675689,0.0017450883,0.0021431737,0.0025673916],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036624891,0.004869542,0.009009195,0.0012678009,0.00043567343,0.0002977168,0.0006396157,0.23532714,0.00547026,0.010417241,0.030344041,0.69825923],"study_design_scores_gemma":[0.00079311774,0.00033497074,0.0012212031,0.000040586932,0.00007632814,0.00008318237,0.00017591362,0.9834501,0.0038086376,0.0072826366,0.0027030783,0.000030279056],"about_ca_topic_score_codex":0.011264586,"about_ca_topic_score_gemma":0.008168401,"teacher_disagreement_score":0.011264586,"about_ca_system_score_codex":0.0014341173,"about_ca_system_score_gemma":0.0020287074,"threshold_uncertainty_score":0.03196639},"labels":[],"label_agreement":null},{"id":"W2950830772","doi":"10.48550/arxiv.1406.2710","title":"A Multiplicative Model for Learning Distributed Text-Based Attribute\\n Representations","year":2014,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research","funders":"","keywords":"Multiplicative function; Computer science; Artificial intelligence; Natural language processing; Theoretical computer science; Mathematics","score_opus":0.13534663944780603,"score_gpt":0.23563582201938707,"score_spread":0.10028918257158104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950830772","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013385292,0.00018059237,0.98477477,0.0004366095,0.000057773217,0.000048845137,0.00023940695,0.0003713556,0.0005053418],"genre_scores_gemma":[0.67103493,0.0007377889,0.310442,0.0007237852,0.0005210838,0.0007241842,0.0021638775,0.00026263483,0.013389698],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9986084,0.00046300143,0.00009393069,0.000499771,0.00022330983,0.00011162772],"domain_scores_gemma":[0.9954992,0.0030696956,0.0003419883,0.0005570454,0.00039018874,0.00014180125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026525636,0.0009844684,0.0013422302,0.0014847295,0.00052569906,0.0015380584,0.0034607288,0.0019248016,0.0028527945],"category_scores_gemma":[0.011151836,0.00075796724,0.0013779288,0.001997225,0.0013632067,0.0046657315,0.0018130575,0.0025451705,0.0011879383],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000383286,0.00031764305,0.0042258175,0.00031144518,0.00024586677,0.00025837077,0.00050936907,0.54470015,0.007855549,0.20689702,0.006837486,0.22745803],"study_design_scores_gemma":[0.000017657432,0.000029166704,0.00018566815,0.000007546961,0.000017026781,0.000028810578,0.0000110534065,0.94819903,0.0006372934,0.050147153,0.00070832396,0.0000112046455],"about_ca_topic_score_codex":0.0023855746,"about_ca_topic_score_gemma":0.0033308733,"teacher_disagreement_score":0.0034607288,"about_ca_system_score_codex":0.0013179114,"about_ca_system_score_gemma":0.0008766493,"threshold_uncertainty_score":0.014028311},"labels":[],"label_agreement":null},{"id":"W2950832322","doi":"10.1007/978-3-662-45912-6_13","title":"Looking at Vector Space and Language Models for IR Using Density Matrices","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Density matrix; Space (punctuation); Vector space; Theoretical computer science; Vector space model; Work (physics); Matrix (chemical analysis); Quantum; Joint (building); Matrix analysis; Algebra over a field; Algorithm; Artificial intelligence; Mathematics; Pure mathematics; Physics; Quantum mechanics","score_opus":0.02338043740809482,"score_gpt":0.25326273663821425,"score_spread":0.22988229923011944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950832322","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066564796,0.0036027816,0.9795694,0.0033074021,0.00022583171,0.000026370704,0.00025111454,0.00063975883,0.0057209954],"genre_scores_gemma":[0.42949176,0.009622948,0.5095029,0.0025045301,0.0020955927,0.0003379209,0.0018655803,0.0012465945,0.043332256],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9972994,0.0016715345,0.00011652928,0.0003697031,0.0003970315,0.00014584896],"domain_scores_gemma":[0.9900034,0.007769228,0.00032825934,0.001038399,0.00065168925,0.0002089895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004313299,0.0011484344,0.0016403157,0.001283863,0.000982444,0.0040984927,0.0022508095,0.002214148,0.010448498],"category_scores_gemma":[0.018179152,0.0011417603,0.002014365,0.0022527967,0.0019842177,0.012437975,0.0018012525,0.0053576827,0.003814927],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000097137985,0.000065817425,0.0007354171,0.00032072476,0.00011183179,0.00008337942,0.00044591862,0.037005648,0.0014828607,0.84967536,0.011300764,0.09867508],"study_design_scores_gemma":[0.000009550823,0.000033523454,0.0001889155,0.000051231607,0.000019935753,0.00008151745,0.00007385235,0.1361949,0.0004715198,0.8580778,0.004765396,0.00003182727],"about_ca_topic_score_codex":0.005057135,"about_ca_topic_score_gemma":0.003727394,"teacher_disagreement_score":0.010448498,"about_ca_system_score_codex":0.0015152019,"about_ca_system_score_gemma":0.0010234022,"threshold_uncertainty_score":0.034953654},"labels":[],"label_agreement":null},{"id":"W2950851846","doi":"10.18653/v1/n19-1323","title":"Connecting Language and Knowledge with Heterogeneous Representations for Neural Relation Extraction","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Relation (database); Computer science; Relationship extraction; Computational linguistics; Artificial intelligence; Natural language processing; Association (psychology); Artificial neural network; Linguistics; Volume (thermodynamics); Knowledge extraction; Cognitive science; Information extraction; Psychology; Epistemology; Philosophy; Data mining","score_opus":0.042728563969961664,"score_gpt":0.3311605949007924,"score_spread":0.28843203093083075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950851846","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07643474,0.00672457,0.89227325,0.0031989033,0.00041382932,0.00023356995,0.003328694,0.009731015,0.007661515],"genre_scores_gemma":[0.65982294,0.0027165571,0.31681186,0.00066041836,0.00051651266,0.00034631888,0.010337584,0.00046525875,0.008322429],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99895346,0.00030062647,0.000100667385,0.00036782867,0.00014856223,0.00012887521],"domain_scores_gemma":[0.9976888,0.0013824626,0.0001237251,0.00047088612,0.0002489045,0.00008523906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014086592,0.0013213684,0.0012001438,0.0044142,0.0008279525,0.002605473,0.002301667,0.0017911496,0.0048303287],"category_scores_gemma":[0.0069053844,0.0008327148,0.0019902003,0.005267059,0.00081716233,0.0075964276,0.0031038017,0.0028040984,0.0029138334],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051110936,0.0003988327,0.0025510606,0.00031949382,0.0003280696,0.0003945841,0.0004560587,0.028651057,0.009718217,0.0172048,0.026111085,0.91335577],"study_design_scores_gemma":[0.000065133114,0.000100565645,0.0013669854,0.000109891014,0.00027814053,0.00021403425,0.00033366223,0.828061,0.008692292,0.15156159,0.0091639105,0.00005283526],"about_ca_topic_score_codex":0.0064103566,"about_ca_topic_score_gemma":0.010164571,"teacher_disagreement_score":0.0064103566,"about_ca_system_score_codex":0.0011833992,"about_ca_system_score_gemma":0.0010222747,"threshold_uncertainty_score":0.016159058},"labels":[],"label_agreement":null},{"id":"W2950912838","doi":"","title":"Enhancing Sentence Embedding with Generalized Pooling","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Pooling; Embedding; Sentence; Computer science; Artificial intelligence; Natural language processing; Inference; Representation (politics); Redundancy (engineering); Theoretical computer science; Machine learning","score_opus":0.06293993825929173,"score_gpt":0.20072024072333788,"score_spread":0.13778030246404616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950912838","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040154457,0.0011315676,0.9533756,0.00034558174,0.00015361971,0.000059670834,0.0002803511,0.002213203,0.0022858535],"genre_scores_gemma":[0.7145345,0.0011978182,0.2720635,0.0005417511,0.00040799522,0.00018322055,0.0014211466,0.0005151808,0.009134848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99955004,0.00015644263,0.00002609178,0.00012745579,0.00008978622,0.00005009791],"domain_scores_gemma":[0.99932885,0.00027398573,0.000065401335,0.0001460743,0.00015337534,0.000032407486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009553122,0.0012064765,0.0007289658,0.00072778267,0.00022379309,0.0006745411,0.00077087333,0.00062181166,0.0024336432],"category_scores_gemma":[0.002812209,0.00023756821,0.0008546127,0.0008430282,0.00034921232,0.0023446633,0.0013156559,0.0008689757,0.0010237211],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027383186,0.00020465314,0.0016509781,0.00040589194,0.00021377036,0.00021345071,0.00048358427,0.08238633,0.08395319,0.016814437,0.01671325,0.7966866],"study_design_scores_gemma":[0.000020932972,0.00020349615,0.0017037761,0.000025195932,0.000118666336,0.00013011288,0.000065346794,0.9427672,0.024895353,0.023425533,0.0066084247,0.0000358953],"about_ca_topic_score_codex":0.0017183144,"about_ca_topic_score_gemma":0.0028698687,"teacher_disagreement_score":0.0024336432,"about_ca_system_score_codex":0.00037012997,"about_ca_system_score_gemma":0.0004162206,"threshold_uncertainty_score":0.008141398},"labels":[],"label_agreement":null},{"id":"W2950935477","doi":"10.18653/v1/p19-1550","title":"Rationally Reappraising ATIS-based Dialogue Systems","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Benchmark (surveying); Grammar; Annotation; Variety (cybernetics); Natural language processing; Set (abstract data type); Artificial intelligence; Task (project management); Domain (mathematical analysis); Process (computing); Information system; Service (business); Programming language; Linguistics","score_opus":0.02052793373822726,"score_gpt":0.23963745582753637,"score_spread":0.21910952208930912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950935477","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.075055465,0.00069385325,0.86049455,0.0017197821,0.00048848806,0.00051796454,0.0012656798,0.04150523,0.018258976],"genre_scores_gemma":[0.47118974,0.00027570734,0.51402605,0.0004510744,0.00015443996,0.00027555932,0.0029183426,0.0029801505,0.007728894],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9924736,0.0037001378,0.0005259004,0.0014051924,0.0015450984,0.0003500887],"domain_scores_gemma":[0.9924281,0.0026040932,0.00032281748,0.0019431008,0.002443089,0.00025874455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056570643,0.0009603034,0.00072524,0.0012183627,0.0008499791,0.003272096,0.0020986702,0.0015469372,0.004257396],"category_scores_gemma":[0.020177677,0.000729374,0.00069100596,0.00058053364,0.0015583223,0.0043033888,0.0042686197,0.00218046,0.0043398985],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056573603,0.00046612,0.0071306564,0.00094699336,0.00012361014,0.00070643605,0.009076577,0.06188284,0.10625159,0.07095928,0.03288214,0.70900804],"study_design_scores_gemma":[0.00008137205,0.00030921111,0.0022768162,0.00020306796,0.00011670623,0.0004265329,0.002416667,0.70583695,0.06204316,0.05896165,0.16718952,0.00013836837],"about_ca_topic_score_codex":0.0049991147,"about_ca_topic_score_gemma":0.0067356555,"teacher_disagreement_score":0.0056570643,"about_ca_system_score_codex":0.0012677478,"about_ca_system_score_gemma":0.0026989062,"threshold_uncertainty_score":0.029917777},"labels":[],"label_agreement":null},{"id":"W2951108824","doi":"","title":"Gated Orthogonal Recurrent Units: On Learning to Forget","year":2017,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Recurrent neural network; Computer science; Forgetting; Treebank; Artificial intelligence; Benchmark (surveying); TIMIT; Natural language processing; Parsing; Speech recognition; Artificial neural network; Hidden Markov model; Cognitive psychology","score_opus":0.13709481040622154,"score_gpt":0.2172512389434832,"score_spread":0.08015642853726165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951108824","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03050353,0.0023068765,0.9566413,0.001261719,0.00032508073,0.000041449886,0.0002786597,0.0017062748,0.006934963],"genre_scores_gemma":[0.7969809,0.002558642,0.17332508,0.0007571685,0.00039386548,0.00021606532,0.0009443306,0.00045139506,0.024372576],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996152,0.00010212913,0.000026093183,0.00013088703,0.000082454746,0.00004316669],"domain_scores_gemma":[0.99910945,0.00042241733,0.0000937526,0.00019839194,0.00012764588,0.00004832575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008053688,0.0010088034,0.00084529247,0.00039666196,0.00029082858,0.00088553864,0.0017448498,0.0010225874,0.0048291786],"category_scores_gemma":[0.0039557493,0.00053768494,0.0006217228,0.0005494383,0.001024024,0.0024252627,0.0014216162,0.0018478612,0.0013170758],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026412052,0.00008313419,0.0013954823,0.00020254048,0.0001256472,0.00025493655,0.00015215554,0.5099719,0.009012284,0.120124765,0.0102178175,0.34819514],"study_design_scores_gemma":[0.000007784869,0.000036084726,0.00011172563,0.000011441478,0.000013436867,0.00002869784,0.000004433359,0.9621861,0.000998111,0.034532413,0.0020625098,0.0000072046632],"about_ca_topic_score_codex":0.0040125186,"about_ca_topic_score_gemma":0.004196561,"teacher_disagreement_score":0.0048291786,"about_ca_system_score_codex":0.00064906356,"about_ca_system_score_gemma":0.00061023684,"threshold_uncertainty_score":0.016155183},"labels":[],"label_agreement":null},{"id":"W2951586423","doi":"10.1111/coin.12225","title":"Multi‐representational convolutional neural networks for text classification","year":2019,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nexen (Canada)","funders":"Natural Science Foundation of Tianjin City; National Natural Science Foundation of China","keywords":"Computer science; Convolutional neural network; Artificial intelligence; Natural language processing; Embedding; Focus (optics); Word embedding; Word (group theory); Semantics (computer science); Categorization; Domain (mathematical analysis); Pattern recognition (psychology)","score_opus":0.0945598607666244,"score_gpt":0.3378425082975396,"score_spread":0.2432826475309152,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951586423","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22983335,0.010238139,0.73504466,0.0027963873,0.0006411823,0.00020995838,0.0027551844,0.008883164,0.009598008],"genre_scores_gemma":[0.8691782,0.0015682864,0.118247576,0.0002609438,0.00021729771,0.00009707623,0.0026205499,0.00009995746,0.00771011],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996712,0.000070590046,0.000027791859,0.00009898145,0.00007107262,0.00006037519],"domain_scores_gemma":[0.99937445,0.00027159462,0.00010117449,0.000077028526,0.0001480282,0.000027734422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006563306,0.0008930746,0.00046538073,0.001490021,0.00022820599,0.000794448,0.0009096417,0.0008251667,0.002368678],"category_scores_gemma":[0.0018397394,0.00023149895,0.00059832423,0.001569522,0.00027190562,0.0015978013,0.0005607789,0.001041569,0.0011932896],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027772484,0.00024115718,0.002979294,0.00026576768,0.00016605103,0.00015097577,0.000089901485,0.27986777,0.014918017,0.007018047,0.011580707,0.6824445],"study_design_scores_gemma":[0.0000027832164,0.00001090581,0.00037104002,0.0000076184574,0.000008256715,0.000008239033,0.0000068219874,0.99534106,0.0014533086,0.0021537286,0.00063271966,0.0000035884198],"about_ca_topic_score_codex":0.008523537,"about_ca_topic_score_gemma":0.008672189,"teacher_disagreement_score":0.008523537,"about_ca_system_score_codex":0.0010693392,"about_ca_system_score_gemma":0.00055230386,"threshold_uncertainty_score":0.016947806},"labels":[],"label_agreement":null},{"id":"W2951626663","doi":"10.48550/arxiv.1809.05524","title":"Extending Neural Generative Conversational Model using External Knowledge Sources","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Perplexity; Computer science; Utterance; Generative grammar; Coherence (philosophical gambling strategy); Natural language processing; Connectionism; Knowledge base; Generative model; Artificial intelligence; Sequence (biology); Focus (optics); Artificial neural network; Language model","score_opus":0.1666798095510348,"score_gpt":0.23416110339318352,"score_spread":0.06748129384214871,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951626663","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.104910396,0.00083799666,0.8811733,0.0009938255,0.00018553482,0.00012418868,0.00054387643,0.0031803185,0.008050446],"genre_scores_gemma":[0.87885344,0.00034762672,0.10976959,0.0003624428,0.00013640821,0.00022222176,0.0014594096,0.00031375128,0.008535093],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99945396,0.00024333176,0.000021092063,0.00017410563,0.000063187275,0.000044409662],"domain_scores_gemma":[0.9982601,0.0012277155,0.000057262787,0.000162932,0.00022122328,0.000070847935],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012788811,0.00091500895,0.0006399997,0.0006912032,0.0005788794,0.0010362664,0.0013954698,0.0010655462,0.0031429764],"category_scores_gemma":[0.004198562,0.00057395897,0.00087109505,0.0005321905,0.00052226265,0.0018311162,0.0013668233,0.0018320804,0.001277197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026477434,0.00024347576,0.0030024215,0.00016323572,0.00020349276,0.0003723762,0.0011020639,0.80479467,0.006804434,0.014943504,0.0054365457,0.16266908],"study_design_scores_gemma":[0.0000048316165,0.000008895198,0.00008466907,0.000004047274,0.000009160983,0.000012740143,0.000011496222,0.9959687,0.00037025881,0.003074537,0.00044667872,0.00000395226],"about_ca_topic_score_codex":0.00962614,"about_ca_topic_score_gemma":0.014091709,"teacher_disagreement_score":0.00962614,"about_ca_system_score_codex":0.0007791633,"about_ca_system_score_gemma":0.0009330515,"threshold_uncertainty_score":0.019140244},"labels":[],"label_agreement":null},{"id":"W2951864354","doi":"10.18653/v1/p19-1423","title":"Inter-sentence Relation Extraction with Document-level Graph Convolutional Neural Network","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"Biotechnology and Biological Sciences Research Council; Associazione Italiana per la Ricerca sul Cancro; National Institute of Advanced Industrial Science and Technology","keywords":"Computer science; Relationship extraction; Sentence; Pairwise comparison; Graph; Artificial intelligence; Convolutional neural network; Exploit; Natural language processing; Relation (database); Theoretical computer science; Information extraction; Data mining","score_opus":0.040599597078830325,"score_gpt":0.26828990305913836,"score_spread":0.22769030598030804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951864354","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08300694,0.0028900378,0.8846659,0.0006101045,0.0001934122,0.00026204032,0.005630214,0.01694561,0.005795791],"genre_scores_gemma":[0.5213899,0.0013482543,0.4382921,0.00037021612,0.00013091696,0.0002876192,0.025948647,0.00071036763,0.011522045],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999617,0.000054898057,0.000028389593,0.00016360344,0.00008977263,0.00004640544],"domain_scores_gemma":[0.9994993,0.00016334381,0.00007617216,0.00010831699,0.00012863624,0.000024131154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004233973,0.0012753748,0.0005861502,0.0021286306,0.00044835388,0.0007522452,0.001094216,0.0009646992,0.0016431796],"category_scores_gemma":[0.0013149343,0.00041576935,0.0010422568,0.002583069,0.00029619902,0.0019040421,0.00078204763,0.0012182025,0.0016094199],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003931814,0.00030374769,0.004874401,0.0005542611,0.00033930418,0.0007940211,0.0003399831,0.09011935,0.0609474,0.011361274,0.03579934,0.79417384],"study_design_scores_gemma":[0.00002056172,0.000058191818,0.0034273686,0.000033739594,0.00011697467,0.00015298142,0.000057519897,0.9511562,0.019303683,0.015800929,0.009841054,0.0000307332],"about_ca_topic_score_codex":0.013093701,"about_ca_topic_score_gemma":0.029221099,"teacher_disagreement_score":0.013093701,"about_ca_system_score_codex":0.0010898536,"about_ca_system_score_gemma":0.0010792362,"threshold_uncertainty_score":0.02603501},"labels":[],"label_agreement":null},{"id":"W2952100657","doi":"10.18653/v1/p19-2027","title":"Dialogue-Act Prediction of Future Responses Based on Conversation History","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Japan Society for the Promotion of Science; Microsoft Research Asia; Microsoft Research","keywords":"Computer science; Utterance; Conversation; Focus (optics); Sequence (biology); Process (computing); Baseline (sea); Artificial intelligence; Chatbot; Recurrent neural network; Speech recognition; Artificial neural network; Machine learning; Natural language processing; Psychology","score_opus":0.022209820183520124,"score_gpt":0.20756981879861972,"score_spread":0.18535999861509958,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952100657","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32586294,0.0018475169,0.65220547,0.0010475399,0.00037716972,0.00042417337,0.0026905565,0.0079284,0.0076162955],"genre_scores_gemma":[0.9150578,0.0003551266,0.07641597,0.0000957921,0.00011664231,0.00022833579,0.0021159924,0.00020949612,0.0054050228],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991579,0.00033168044,0.000043619344,0.00028092763,0.00012458458,0.000061284],"domain_scores_gemma":[0.99675995,0.0020790477,0.00019910636,0.00015442653,0.0006275634,0.0001799152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015092618,0.0012442073,0.00059903465,0.0015024194,0.00031271536,0.0007315138,0.0006606032,0.0006875618,0.0027858035],"category_scores_gemma":[0.005000927,0.00034693378,0.000744419,0.0004397094,0.00028301508,0.0010629356,0.00058118056,0.001300815,0.0024365925],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024154016,0.0007319171,0.044628017,0.0010653269,0.00040289146,0.00067970075,0.002299863,0.22325888,0.064278945,0.0058903,0.017505767,0.63684297],"study_design_scores_gemma":[0.000016277167,0.00011700732,0.005653833,0.000029894043,0.000049160994,0.0000761037,0.00014342861,0.983712,0.006346335,0.0018142327,0.0020179416,0.000023737997],"about_ca_topic_score_codex":0.004744263,"about_ca_topic_score_gemma":0.0058138734,"teacher_disagreement_score":0.004744263,"about_ca_system_score_codex":0.00044641344,"about_ca_system_score_gemma":0.00063016114,"threshold_uncertainty_score":0.0094332695},"labels":[],"label_agreement":null},{"id":"W2952243641","doi":"10.4000/books.aaccademia.4638","title":"UO_IRO: Linguistic informed deep-learning model for irony detection","year":2018,"lang":"en","type":"book-chapter","venue":"Accademia University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Research Council Canada; Università degli Studi di Napoli Federico II","keywords":"Task (project management); Convolutional neural network; Computer science; Artificial intelligence; Deep learning; Irony; Linguistics; Natural language processing; Psychology; Philosophy; Engineering","score_opus":0.03661978327563584,"score_gpt":0.23238101320675986,"score_spread":0.19576122993112402,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952243641","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0354961,0.0044978745,0.839516,0.0015211181,0.0012863188,0.00031423848,0.0065578013,0.065866895,0.044943705],"genre_scores_gemma":[0.2600234,0.0021810795,0.63225645,0.0009089521,0.00044729718,0.00046023843,0.02156705,0.0035875847,0.07856797],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997757,0.00003990444,0.000011419052,0.00008463842,0.000056216406,0.000032085565],"domain_scores_gemma":[0.99977714,0.000080356476,0.00001743473,0.000042668336,0.000061588304,0.000020835063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005634951,0.0014714741,0.0007507277,0.0010357946,0.000501015,0.0012772323,0.0012437517,0.0011322598,0.009955623],"category_scores_gemma":[0.0012356308,0.00050754705,0.00084797974,0.0006141608,0.0002821847,0.0021412026,0.0012298872,0.0015952246,0.0066085993],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003368938,0.00022043174,0.0021144645,0.0005327913,0.00010641506,0.00032309262,0.0004187706,0.024916856,0.022445733,0.015076292,0.18113558,0.7523727],"study_design_scores_gemma":[0.000047542402,0.00011627374,0.0022761594,0.00018787557,0.00008768603,0.00039699362,0.00021920934,0.78015625,0.022419091,0.029362386,0.164648,0.00008261824],"about_ca_topic_score_codex":0.0034866303,"about_ca_topic_score_gemma":0.0084087225,"teacher_disagreement_score":0.009955623,"about_ca_system_score_codex":0.0006564648,"about_ca_system_score_gemma":0.0007763984,"threshold_uncertainty_score":0.03330487},"labels":[],"label_agreement":null},{"id":"W2952268267","doi":"10.48550/arxiv.1810.01375","title":"A Knowledge Hunting Framework for Common Sense Reasoning","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Schema (genetic algorithms); Computer science; Inference; Task (project management); Common sense; Commonsense reasoning; Inference engine; Artificial intelligence; Machine learning; Epistemology; Engineering","score_opus":0.10324345316713325,"score_gpt":0.2338531365749794,"score_spread":0.13060968340784612,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952268267","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002053759,0.00030004166,0.9894288,0.0005498538,0.000058629645,0.00017143402,0.0004287729,0.00514943,0.0018593385],"genre_scores_gemma":[0.05936823,0.00023524027,0.93554735,0.00037471383,0.00011484853,0.00026419404,0.001859874,0.0004351416,0.0018003597],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99287146,0.002235641,0.0005998768,0.0019586757,0.0020599165,0.00027445838],"domain_scores_gemma":[0.9896685,0.005449838,0.00051619456,0.002967983,0.0010427726,0.0003547053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007865506,0.0016734009,0.0014700398,0.007944485,0.0021800802,0.005185439,0.0063733044,0.002584706,0.009453135],"category_scores_gemma":[0.022514261,0.0012027989,0.0037268165,0.004300797,0.0029521622,0.011876562,0.0075233094,0.004761516,0.0042772014],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025731954,0.0005178315,0.002788752,0.0009494868,0.00037040445,0.0005923812,0.0029844793,0.02362715,0.012213438,0.25366616,0.039346643,0.662686],"study_design_scores_gemma":[0.000053991535,0.00006124604,0.0006222646,0.00015858829,0.00008922846,0.0005158898,0.00047535758,0.49225354,0.009381227,0.4508141,0.045488857,0.00008567198],"about_ca_topic_score_codex":0.006227805,"about_ca_topic_score_gemma":0.008700919,"teacher_disagreement_score":0.009453135,"about_ca_system_score_codex":0.0015540821,"about_ca_system_score_gemma":0.0024707022,"threshold_uncertainty_score":0.041597247},"labels":[],"label_agreement":null},{"id":"W2952349219","doi":"10.18653/v1/p19-1166","title":"Understanding Undesirable Word Embedding Associations","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":121,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Word (group theory); Debiasing; Association (psychology); Subspace topology; Computer science; Embedding; Projection (relational algebra); Natural language processing; Product (mathematics); Artificial intelligence; Association test; Word Association; Word embedding; Matrix decomposition; Speech recognition; Algorithm; Mathematics; Psychology; Social psychology","score_opus":0.1460175943364272,"score_gpt":0.29422688950393994,"score_spread":0.14820929516751274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952349219","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34265152,0.0008186165,0.6495171,0.0011116505,0.00020860518,0.00006656774,0.0006687383,0.0006881474,0.004269043],"genre_scores_gemma":[0.9360655,0.0003229344,0.060863957,0.00027840433,0.00013379738,0.00009861207,0.0008437251,0.00020619112,0.0011869593],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99229926,0.003561903,0.00064620306,0.0016635458,0.0014656631,0.00036349564],"domain_scores_gemma":[0.9631541,0.024492027,0.0037051463,0.004224978,0.003962374,0.000461401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074969972,0.0007237821,0.0007916691,0.0016510843,0.0009876014,0.002501348,0.000769073,0.0012180976,0.0020761003],"category_scores_gemma":[0.06115572,0.0005182474,0.00044756345,0.0016370144,0.0018229228,0.006437663,0.0037464346,0.0020588082,0.00074792566],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008557453,0.0002720597,0.14909889,0.0007598457,0.0003772876,0.000892394,0.007694384,0.045646116,0.029628674,0.29095632,0.013068685,0.4607496],"study_design_scores_gemma":[0.000047544298,0.00023822475,0.025175368,0.00015678811,0.00012198427,0.0014234176,0.002453208,0.46371788,0.020018859,0.46837655,0.01817102,0.000099199744],"about_ca_topic_score_codex":0.0010119091,"about_ca_topic_score_gemma":0.0015941081,"teacher_disagreement_score":0.0074969972,"about_ca_system_score_codex":0.0005294199,"about_ca_system_score_gemma":0.0008697316,"threshold_uncertainty_score":0.039648414},"labels":[],"label_agreement":null},{"id":"W2952383053","doi":"10.18653/v1/p19-1145","title":"Lightweight and Efficient Neural Natural Language Processing with Quaternion Networks","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Quaternion; Hypercomplex number; Computer science; Parameterized complexity; Artificial neural network; Artificial intelligence; Computation; Quaternion algebra; Theoretical computer science; Algorithm; Algebra over a field; Mathematics; Pure mathematics","score_opus":0.008538991976912488,"score_gpt":0.2300345345057534,"score_spread":0.2214955425288409,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952383053","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022880374,0.00037998427,0.96986604,0.00034239652,0.00007876159,0.000046881,0.0001568155,0.0030530775,0.003195705],"genre_scores_gemma":[0.669286,0.00077654846,0.3209471,0.0003102172,0.00011471451,0.00020282589,0.000648645,0.0003195019,0.0073944284],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998223,0.0000382528,0.000012976713,0.000049061815,0.00005720922,0.000020187015],"domain_scores_gemma":[0.9996685,0.00011477507,0.000038647973,0.00009989106,0.000060316954,0.00001794886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00035602346,0.00059644476,0.00040747988,0.00032861275,0.00025464222,0.00084728154,0.0014314197,0.00054728653,0.0046709185],"category_scores_gemma":[0.0013851699,0.0003076768,0.00048475157,0.0005305247,0.00045920524,0.002380238,0.0009740323,0.0010000584,0.001624735],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027694958,0.00012799528,0.00075393345,0.00026049605,0.00011519526,0.00019988052,0.00021645715,0.44373256,0.058681984,0.07179574,0.009060192,0.41477862],"study_design_scores_gemma":[0.00000979685,0.000027089081,0.00008252864,0.000005894387,0.000008757929,0.000023360044,0.000012041567,0.97086203,0.0065830243,0.019996762,0.0023815138,0.0000073415545],"about_ca_topic_score_codex":0.0033075132,"about_ca_topic_score_gemma":0.0050920215,"teacher_disagreement_score":0.0046709185,"about_ca_system_score_codex":0.0006494187,"about_ca_system_score_gemma":0.00054417306,"threshold_uncertainty_score":0.015625775},"labels":[],"label_agreement":null},{"id":"W2952802110","doi":"10.18653/v1/p19-1030","title":"You Only Need Attention to Traverse Trees","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Tree traversal; Sentence; Traverse; Artificial intelligence; Phrase; Tree structure; Natural language processing; Dependency (UML); Tree (set theory); Dependency grammar; Binary tree; Algorithm; Mathematics","score_opus":0.016250234093927878,"score_gpt":0.2314662255295989,"score_spread":0.21521599143567102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952802110","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022345098,0.00041709814,0.96476835,0.0013638559,0.00011094484,0.00006408062,0.0004178799,0.0027993314,0.0077132927],"genre_scores_gemma":[0.5507485,0.0013197047,0.42332822,0.0009931797,0.00013453716,0.00018755408,0.0014306044,0.0008762155,0.020981492],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99977976,0.000050392217,0.00001033542,0.00009082951,0.000038734972,0.000029954967],"domain_scores_gemma":[0.9993351,0.00032367892,0.000040955947,0.00016609226,0.00009950668,0.000034613317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004902149,0.0005009147,0.00045017374,0.00034632804,0.00039159672,0.0009636543,0.0008958312,0.0008768135,0.0071271122],"category_scores_gemma":[0.0029713248,0.00045929715,0.00060058385,0.00054226967,0.0005204265,0.0048652096,0.0009455564,0.0015288383,0.0022598042],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028683734,0.00010174229,0.0022917003,0.0005051769,0.000110574576,0.00028561393,0.0010965242,0.07920289,0.04489047,0.18679674,0.024617657,0.6598141],"study_design_scores_gemma":[0.000023212076,0.00010806778,0.0013053477,0.00007538097,0.000073561285,0.00026303795,0.00016007798,0.6436318,0.014346283,0.30552256,0.034449235,0.00004148847],"about_ca_topic_score_codex":0.0040076505,"about_ca_topic_score_gemma":0.0073195696,"teacher_disagreement_score":0.0071271122,"about_ca_system_score_codex":0.0005265634,"about_ca_system_score_gemma":0.0006303329,"threshold_uncertainty_score":0.023842573},"labels":[],"label_agreement":null},{"id":"W2952879766","doi":"10.1142/9789811203527_0003","title":"Natural Language Processing: An Overview","year":2019,"lang":"en","type":"book-chapter","venue":"WORLD SCIENTIFIC eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Natural language processing","score_opus":0.054599631079788005,"score_gpt":0.2866057392801148,"score_spread":0.2320061082003268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952879766","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010917654,0.6922554,0.16389495,0.0031245747,0.0026893702,0.0002662342,0.0014621393,0.0041012326,0.13111442],"genre_scores_gemma":[0.009086314,0.69337535,0.14153285,0.0025334437,0.0049719494,0.00054328307,0.005289726,0.0019110892,0.14075603],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993168,0.0000910732,0.00006508968,0.00014821034,0.00033583626,0.00004305498],"domain_scores_gemma":[0.999008,0.00057157886,0.000033856824,0.00008038506,0.00025783197,0.000048381706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001079701,0.0014467459,0.0011074872,0.0064395117,0.0006220716,0.0046709315,0.0016992997,0.0016312561,0.022228895],"category_scores_gemma":[0.0019039243,0.0011438422,0.0010696617,0.010200842,0.0009061085,0.0075410255,0.001690318,0.0025589971,0.024003265],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020666681,0.00006890142,0.00008976555,0.002236332,0.000026864409,0.00006648478,0.00016514726,0.0010361188,0.0018176452,0.023661196,0.14235505,0.8284558],"study_design_scores_gemma":[0.000005878713,0.000020733349,0.00031268902,0.0005802699,0.000019771169,0.00042882827,0.000068888876,0.0017969738,0.0009832626,0.029909551,0.96584845,0.000024684823],"about_ca_topic_score_codex":0.0024865617,"about_ca_topic_score_gemma":0.0033507547,"teacher_disagreement_score":0.022228895,"about_ca_system_score_codex":0.0012343046,"about_ca_system_score_gemma":0.0015765948,"threshold_uncertainty_score":0.07436305},"labels":[],"label_agreement":null},{"id":"W2952944718","doi":"10.18653/v1/w19-5004","title":"REflex: Flexible Framework for Relation Extraction in Multiple Domains","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; Mitacs","keywords":"Computer science; Relation (database); Field (mathematics); Data science; Artificial intelligence; Data mining; Machine learning; Mathematics","score_opus":0.08271429853203008,"score_gpt":0.35071644763255033,"score_spread":0.2680021491005202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952944718","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017909658,0.0011175629,0.93712074,0.00061622023,0.00011684562,0.00027291317,0.0095971115,0.046549324,0.0028183025],"genre_scores_gemma":[0.04392413,0.00090198946,0.91980827,0.0004569493,0.0001428301,0.0005041125,0.025073042,0.002838271,0.00635036],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9968106,0.0008471822,0.00031636693,0.0012140238,0.00063332944,0.00017842544],"domain_scores_gemma":[0.9965145,0.0016289275,0.00021849874,0.0011841019,0.00029888432,0.00015515603],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045313,0.0018884499,0.0012328981,0.0072094365,0.0011296053,0.0042113084,0.0028032481,0.002265197,0.02023111],"category_scores_gemma":[0.011826549,0.0012035455,0.0038894815,0.0050048083,0.00081042945,0.0077777724,0.005335672,0.003368786,0.019435136],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036901652,0.00022676069,0.0032297196,0.001969255,0.00040182014,0.0006525302,0.0009556994,0.015764236,0.009473379,0.082314014,0.17589645,0.708747],"study_design_scores_gemma":[0.00014804566,0.00011904779,0.003065169,0.0006034151,0.0002201196,0.0014738167,0.00082957867,0.2959286,0.014222043,0.24390154,0.4393501,0.00013860465],"about_ca_topic_score_codex":0.0047145057,"about_ca_topic_score_gemma":0.012973239,"teacher_disagreement_score":0.02023111,"about_ca_system_score_codex":0.0010377804,"about_ca_system_score_gemma":0.0021631606,"threshold_uncertainty_score":0.06767976},"labels":[],"label_agreement":null},{"id":"W2952989203","doi":"10.18653/v1/p19-1189","title":"Reinforced Training Data Selection for Domain Adaptation","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Task (project management); Artificial intelligence; Machine learning; Selection (genetic algorithm); Domain (mathematical analysis); Reinforcement learning; Process (computing); Adaptation (eye); Representation (politics); Set (abstract data type); Generator (circuit theory); Parsing; Domain adaptation; Stability (learning theory); Data mining","score_opus":0.11472330403549424,"score_gpt":0.2895308221460234,"score_spread":0.17480751811052914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2952989203","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028950179,0.00021715216,0.9675354,0.000207114,0.000048536353,0.00013227464,0.0001072589,0.0018763263,0.0009258706],"genre_scores_gemma":[0.7434579,0.00017480267,0.25210914,0.0003557479,0.00007826588,0.0005342576,0.0006683636,0.00029508388,0.0023264606],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998588,0.00065433525,0.00006978091,0.00042459898,0.0001769218,0.00008640796],"domain_scores_gemma":[0.99625325,0.0020531768,0.0001944616,0.00091041066,0.00045326527,0.0001354616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028287852,0.0009412191,0.0010176755,0.00066324614,0.00037657667,0.0005986771,0.0019199245,0.00092051923,0.0016227198],"category_scores_gemma":[0.00966408,0.0004830297,0.00061903545,0.00063177344,0.00093407446,0.0017608544,0.0019944534,0.0020422058,0.0007803995],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041450027,0.00048914115,0.005768545,0.00017917747,0.0001646136,0.00019511543,0.00035541278,0.43060964,0.02021202,0.013073616,0.0068080258,0.52173024],"study_design_scores_gemma":[0.000027798831,0.0000577926,0.0003723137,0.000008378264,0.000013425096,0.00003899018,0.000021735234,0.9868579,0.004124986,0.0071895947,0.0012761526,0.000010857427],"about_ca_topic_score_codex":0.0016215784,"about_ca_topic_score_gemma":0.0024908017,"teacher_disagreement_score":0.0028287852,"about_ca_system_score_codex":0.000651723,"about_ca_system_score_gemma":0.0011134094,"threshold_uncertainty_score":0.014960289},"labels":[],"label_agreement":null},{"id":"W2953476352","doi":"10.1007/978-3-031-24337-0_24","title":"Multiplicative Models for Recurrent Language Modeling","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Recurrent neural network; Language model; Multiplicative function; Artificial intelligence; Relevance (law); Parametrization (atmospheric modeling); Simple (philosophy); Theoretical computer science; Artificial neural network; Natural language processing; Machine learning; Mathematics","score_opus":0.05669957815851957,"score_gpt":0.29500257596396456,"score_spread":0.238302997805445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953476352","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001219698,0.0015123137,0.9938624,0.00031447396,0.0001948112,0.000012379334,0.00014742107,0.00047875728,0.0022578149],"genre_scores_gemma":[0.2572859,0.009316481,0.64367396,0.0013792528,0.00271003,0.00069129356,0.003316183,0.0024475085,0.07917938],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881124,0.0006008621,0.00007734623,0.00023652443,0.00019663265,0.00007745688],"domain_scores_gemma":[0.9958668,0.0031788733,0.0001468918,0.00046722937,0.00027005814,0.00007023629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001973034,0.0015733584,0.0016034701,0.00078874454,0.0004471524,0.0019889541,0.002615379,0.0019781762,0.009858529],"category_scores_gemma":[0.0086027235,0.0011103384,0.0018313208,0.0013516776,0.0006986104,0.0034904694,0.0015796428,0.003705911,0.006141008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010617702,0.00011495626,0.00045819703,0.00047506028,0.00026585435,0.00021050408,0.0002951361,0.13083403,0.0049736775,0.61545396,0.028179158,0.21863326],"study_design_scores_gemma":[0.0000112352345,0.000026163892,0.000112621885,0.000029910949,0.000048094866,0.000078310324,0.000016424725,0.66614556,0.00074744545,0.32409117,0.008669523,0.000023543898],"about_ca_topic_score_codex":0.001679414,"about_ca_topic_score_gemma":0.0031253907,"teacher_disagreement_score":0.009858529,"about_ca_system_score_codex":0.0007671403,"about_ca_system_score_gemma":0.0006007834,"threshold_uncertainty_score":0.032980025},"labels":[],"label_agreement":null},{"id":"W2953591948","doi":"10.18653/v1/s19-2188","title":"UBC-NLP at SemEval-2019 Task 4: Hyperpartisan News Detection With Attention-Based Bi-LSTMs","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Western Canada Research Grid; Compute Canada","keywords":"SemEval; Computer science; Task (project management); Ranking (information retrieval); Artificial intelligence; Natural language processing; Competition (biology); Machine learning","score_opus":0.010757826205529956,"score_gpt":0.2112432914759302,"score_spread":0.20048546527040026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953591948","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29102248,0.010990013,0.17510916,0.0102871405,0.016480297,0.0026165617,0.20303588,0.21143769,0.07902088],"genre_scores_gemma":[0.3343369,0.0011208417,0.18393318,0.002551558,0.0012706404,0.0016100042,0.39947587,0.007857332,0.06784365],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99558675,0.0012492068,0.00022128069,0.0014924523,0.0008553637,0.00059504376],"domain_scores_gemma":[0.99390996,0.0018371338,0.00019228309,0.001442533,0.0020409338,0.0005772055],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006251506,0.004555582,0.0021574744,0.0024423152,0.0019448511,0.0037335502,0.0030317686,0.0050741285,0.026942993],"category_scores_gemma":[0.013931755,0.0010483708,0.0017947365,0.0023066907,0.0008125679,0.004729112,0.0040107216,0.0048435386,0.03135184],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016081322,0.0009255167,0.0033478704,0.0010085835,0.00045318782,0.00072936056,0.00033675492,0.019591864,0.015045825,0.0018996166,0.6832392,0.27181414],"study_design_scores_gemma":[0.001420893,0.0012761448,0.008904081,0.00034207193,0.00035927922,0.0011905169,0.0009464321,0.6616378,0.06768806,0.011306091,0.24463262,0.0002959042],"about_ca_topic_score_codex":0.01712045,"about_ca_topic_score_gemma":0.02386935,"teacher_disagreement_score":0.026942993,"about_ca_system_score_codex":0.0018456148,"about_ca_system_score_gemma":0.002519982,"threshold_uncertainty_score":0.09013331},"labels":[],"label_agreement":null},{"id":"W2953642885","doi":"10.1145/3331184.3331354","title":"Dynamic Sampling Meets Pooling","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Institute of Standards and Technology","keywords":"Pooling; NIST; Computer science; Sampling (signal processing); Statistics; Set (abstract data type); Sample (material); Data mining; Information retrieval; Artificial intelligence; Natural language processing; Mathematics; Computer vision","score_opus":0.020515901836123764,"score_gpt":0.26266280865947483,"score_spread":0.24214690682335108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953642885","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12044381,0.0012707819,0.8476253,0.0008900033,0.00041337757,0.0021156606,0.0017747359,0.0047593876,0.020706868],"genre_scores_gemma":[0.68157744,0.00027520768,0.30357406,0.00066333555,0.0004135886,0.002353425,0.0026495722,0.001371281,0.0071220747],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97846687,0.009210346,0.001439122,0.0048716874,0.004851156,0.0011608438],"domain_scores_gemma":[0.9651088,0.016927842,0.0013555982,0.00983224,0.0061075473,0.00066797814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02229661,0.0022581117,0.002069941,0.0018201466,0.0019220547,0.004109826,0.0021848339,0.0016598261,0.0076844324],"category_scores_gemma":[0.075523324,0.0008710079,0.0017995224,0.0025341697,0.0014499457,0.006243252,0.0055884016,0.0021029967,0.0027883044],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026405219,0.0007006987,0.023547087,0.001299858,0.0007481939,0.00037025675,0.0049091415,0.056894697,0.039727263,0.028139591,0.03157566,0.8094471],"study_design_scores_gemma":[0.0010400001,0.0033104608,0.057546347,0.00039276402,0.0013207293,0.001518313,0.0037978047,0.5317167,0.090180516,0.17272422,0.13581657,0.0006356849],"about_ca_topic_score_codex":0.006194763,"about_ca_topic_score_gemma":0.006894866,"teacher_disagreement_score":0.02229661,"about_ca_system_score_codex":0.0016883131,"about_ca_system_score_gemma":0.0023541444,"threshold_uncertainty_score":0.11791712},"labels":[],"label_agreement":null},{"id":"W2953702336","doi":"10.18653/v1/w19-2511","title":"Semantics and Homothetic Clustering of Hafez Poetry","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Homothetic transformation; Cluster analysis; Poetry; Computer science; Similarity (geometry); Semantics (computer science); Feature (linguistics); Natural language processing; Artificial intelligence; Word (group theory); Mathematics; Information retrieval; Linguistics; Image (mathematics); Philosophy","score_opus":0.010045959442382551,"score_gpt":0.2146126265124282,"score_spread":0.20456666707004567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953702336","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8264995,0.0002955831,0.14988671,0.00042199064,0.00012208894,0.00019224823,0.0012179698,0.0008544867,0.020509467],"genre_scores_gemma":[0.9241813,0.00008168963,0.07116489,0.000040479812,0.00003654822,0.00010823727,0.0012999322,0.00013944195,0.0029474501],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989581,0.00036756496,0.00007891713,0.00029555353,0.00023604563,0.000063791886],"domain_scores_gemma":[0.9973838,0.0010790116,0.00035063468,0.00045582338,0.0006085977,0.00012207897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012759677,0.00033221257,0.00026263666,0.0045367773,0.0015331417,0.0019680948,0.00035411195,0.0004388999,0.002574343],"category_scores_gemma":[0.006155964,0.00018579072,0.00034003655,0.0026787566,0.0016364363,0.0020021573,0.0008967624,0.00071054476,0.0006909124],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013277796,0.00038488294,0.059147116,0.0007425565,0.00013274308,0.00062399666,0.025919693,0.01975354,0.049920898,0.18436536,0.0109628355,0.6467186],"study_design_scores_gemma":[0.0001951156,0.0006262323,0.25866205,0.00034407174,0.00014518364,0.0022919015,0.02003896,0.31874532,0.0711479,0.20514672,0.122332335,0.00032420398],"about_ca_topic_score_codex":0.0016852362,"about_ca_topic_score_gemma":0.002724024,"teacher_disagreement_score":0.0045367773,"about_ca_system_score_codex":0.0008532688,"about_ca_system_score_gemma":0.00049754133,"threshold_uncertainty_score":0.008612096},"labels":[],"label_agreement":null},{"id":"W2953876178","doi":"10.1609/icwsm.v13i01.3246","title":"PhAITV: A Phrase Author Interaction Topic Viewpoint Model for the Summarization of Reasons Expressed by Polarized Stances","year":2019,"lang":"en","type":"article","venue":"Proceedings of the International AAAI Conference on Web and Social Media","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Automatic summarization; Computer science; Phrase; Viewpoints; Relevance (law); Natural language processing; Cluster analysis; Argument (complex analysis); Pipeline (software); Artificial intelligence; Information retrieval; Linguistics","score_opus":0.04469682486867616,"score_gpt":0.2824244453236807,"score_spread":0.23772762045500456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953876178","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025145857,0.0017914951,0.95379317,0.0006275535,0.00021078042,0.00043867447,0.0068473998,0.007181654,0.0039634695],"genre_scores_gemma":[0.22712883,0.0010787909,0.73518515,0.00021365794,0.0003432643,0.00076235464,0.028667565,0.00092880934,0.00569146],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993025,0.00022970728,0.00005272682,0.0001895388,0.0001770008,0.000048400398],"domain_scores_gemma":[0.9985948,0.00071773306,0.0001426394,0.00019881957,0.00027294864,0.00007301263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012940753,0.0010110096,0.00063108915,0.002857833,0.000662699,0.002031252,0.0013821088,0.001277006,0.0033853047],"category_scores_gemma":[0.004982272,0.00040384385,0.001383655,0.0019092598,0.00038513102,0.0023637265,0.0014420446,0.0017858453,0.0022125188],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067547115,0.0002459884,0.007852371,0.0011455162,0.00036327523,0.00047104651,0.0030033495,0.035407394,0.026030151,0.035597812,0.041889057,0.8473185],"study_design_scores_gemma":[0.0001234881,0.00022912175,0.0069012935,0.00018719037,0.00024857558,0.0003057424,0.0009600155,0.8677611,0.0125584,0.048882645,0.06175829,0.00008417014],"about_ca_topic_score_codex":0.0044982173,"about_ca_topic_score_gemma":0.009890987,"teacher_disagreement_score":0.0044982173,"about_ca_system_score_codex":0.000920721,"about_ca_system_score_gemma":0.0014220042,"threshold_uncertainty_score":0.011325002},"labels":[],"label_agreement":null},{"id":"W2954222332","doi":"10.1016/j.knosys.2019.104925","title":"Context-aware instance matching through graph embedding in lexical semantic space","year":2019,"lang":"en","type":"article","venue":"Knowledge-Based Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario College of Art and Design; Université du Québec à Chicoutimi; Université du Québec à Montréal","funders":"Fonds de recherche du Québec – Nature et technologies; Canada Foundation for Innovation; Ministère de l'Économie, de la Science et de l'Innovation - Québec","keywords":"Computer science; Ontology alignment; Knowledge graph; Embedding; Information retrieval; RDF; Semantic Web; Entity linking; Matching (statistics); Knowledge base; Referent; Theoretical computer science; Artificial intelligence; Ontology-based data integration","score_opus":0.02465723400931598,"score_gpt":0.2784244017600168,"score_spread":0.25376716775070085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2954222332","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08682299,0.0017821912,0.9003023,0.00042590254,0.00017731187,0.00016804916,0.0020378625,0.0048487615,0.0034347146],"genre_scores_gemma":[0.6461223,0.0009081515,0.34503594,0.00014341081,0.000100936275,0.000101714344,0.004715105,0.00044650552,0.0024259316],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99875927,0.0002588102,0.000111956884,0.00046825228,0.00029909654,0.00010271469],"domain_scores_gemma":[0.99897027,0.00042787808,0.000087292734,0.00026592345,0.00019133228,0.00005732468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00053964293,0.0006020192,0.0011593242,0.004326863,0.00069890876,0.002159249,0.0013726288,0.0010938044,0.002641516],"category_scores_gemma":[0.003567167,0.00038916848,0.0012028825,0.005113021,0.00041977657,0.005147661,0.0023078246,0.0009983268,0.0011116454],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008414122,0.0006924745,0.00792426,0.0007441402,0.0004599913,0.000519987,0.000757013,0.043175567,0.033530544,0.04244311,0.014652578,0.854259],"study_design_scores_gemma":[0.00005062845,0.000107896994,0.0027087806,0.00007082626,0.00027969352,0.00037231224,0.0005336953,0.88019735,0.016544024,0.088765055,0.010313337,0.000056353463],"about_ca_topic_score_codex":0.0057228915,"about_ca_topic_score_gemma":0.009042091,"teacher_disagreement_score":0.0057228915,"about_ca_system_score_codex":0.00049491285,"about_ca_system_score_gemma":0.0009972769,"threshold_uncertainty_score":0.011379123},"labels":[],"label_agreement":null},{"id":"W2955144464","doi":"10.18653/v1/w18-5509","title":"Automated Fact-Checking of Claims in Argumentative Parliamentary Debates","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; University of Toronto","keywords":"Argumentative; Leverage (statistics); Computer science; Artificial intelligence; Natural language processing; Parliament; Class (philosophy); Measure (data warehouse); Binary classification; Binary number; Machine learning; Linguistics; Data mining; Support vector machine; Mathematics; Law; Political science; Arithmetic","score_opus":0.02663764514536591,"score_gpt":0.29502840751279213,"score_spread":0.2683907623674262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955144464","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8326627,0.0015679979,0.11763028,0.0030696075,0.00027248787,0.00037801062,0.013416788,0.008912772,0.022089401],"genre_scores_gemma":[0.94227564,0.0001582777,0.044104595,0.00010731239,0.000077261495,0.000051372066,0.010641298,0.00016272151,0.0024216117],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99590236,0.0013259359,0.00027507258,0.00084516685,0.0012577088,0.00039379043],"domain_scores_gemma":[0.96716475,0.019031031,0.005022132,0.0027161366,0.0053633046,0.00070265523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00471008,0.0006494942,0.0005989439,0.006466862,0.0016394773,0.0033183205,0.0014493936,0.0016110934,0.0031163485],"category_scores_gemma":[0.028112674,0.00038631284,0.00067000464,0.0026999358,0.000892349,0.0029926938,0.0020058127,0.002058387,0.001996785],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008716153,0.0004863963,0.2947416,0.00086818734,0.00022959198,0.001309352,0.007886604,0.03259339,0.03200964,0.019692816,0.04167691,0.5676339],"study_design_scores_gemma":[0.000045976656,0.000073275325,0.14419319,0.00029750424,0.00012787453,0.0006641029,0.0033160665,0.7476405,0.04160696,0.022825468,0.03909669,0.000112358546],"about_ca_topic_score_codex":0.03600551,"about_ca_topic_score_gemma":0.06519457,"teacher_disagreement_score":0.03600551,"about_ca_system_score_codex":0.0024344039,"about_ca_system_score_gemma":0.0031116607,"threshold_uncertainty_score":0.071591854},"labels":[],"label_agreement":null},{"id":"W2955236073","doi":"10.18653/v1/w19-5023","title":"Enhancing PIO Element Detection in Medical Text Using Contextualized Embedding","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"International Medias Data Services (Canada)","funders":"","keywords":"Embedding; Computer science; Classifier (UML); Artificial intelligence; Encoder; Transformer; Ambiguity; Boosting (machine learning); Machine learning; Leverage (statistics); Population; Pattern recognition (psychology); Data mining; Engineering; Medicine","score_opus":0.03735605152197291,"score_gpt":0.32744130695466706,"score_spread":0.29008525543269414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2955236073","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1619114,0.0057877945,0.81902564,0.0017903046,0.00049946754,0.00030887127,0.0039048765,0.002336081,0.004435621],"genre_scores_gemma":[0.69096035,0.0014588863,0.29802915,0.00044015938,0.0007389089,0.00022525157,0.0052578845,0.00021264286,0.002676716],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985405,0.0006405952,0.00013949657,0.0003988988,0.00020416119,0.000076339325],"domain_scores_gemma":[0.9906825,0.0065476107,0.0010104313,0.00091797,0.00070104963,0.00014040207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022488334,0.0007084247,0.00060302933,0.0021349725,0.0002328919,0.0009403801,0.00056787534,0.0010362788,0.0014966846],"category_scores_gemma":[0.015002638,0.00020012562,0.00056804233,0.001215366,0.0004246372,0.002130114,0.0009633738,0.001025825,0.0011227743],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008563165,0.0005318702,0.040525433,0.00133444,0.00023689744,0.00042448248,0.0008428424,0.047521062,0.034203544,0.008090582,0.008513836,0.85691875],"study_design_scores_gemma":[0.00008637319,0.00089870766,0.031083915,0.00040730857,0.00032782738,0.0014289909,0.000468266,0.85815054,0.040482108,0.040479243,0.026067166,0.000119515535],"about_ca_topic_score_codex":0.0006487456,"about_ca_topic_score_gemma":0.0013138659,"teacher_disagreement_score":0.0022488334,"about_ca_system_score_codex":0.00029636032,"about_ca_system_score_gemma":0.0004609327,"threshold_uncertainty_score":0.011893153},"labels":[],"label_agreement":null},{"id":"W2962471269","doi":"10.1145/3319619.3321957","title":"Evolving recurrent neural networks for emergent communication","year":2019,"lang":"en","type":"article","venue":"Proceedings of the Genetic and Evolutionary Computation Conference Companion","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Neuroevolution; Computer science; Reinforcement learning; Artificial intelligence; Artificial neural network; Recurrent neural network; Machine learning","score_opus":0.02828320261727017,"score_gpt":0.25084210473312984,"score_spread":0.22255890211585969,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2962471269","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054168813,0.0005722013,0.93371195,0.0006381072,0.00008431643,0.000035508718,0.00007292936,0.00071467547,0.010001553],"genre_scores_gemma":[0.88712734,0.00047004232,0.103707895,0.00019273395,0.000060989732,0.00015525588,0.00014504655,0.00020030612,0.00794046],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99975306,0.00007209086,0.000014140606,0.00006556715,0.00006144086,0.00003355952],"domain_scores_gemma":[0.9994272,0.00029303163,0.00009556281,0.00006246517,0.00008116052,0.000040592822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00066861475,0.00046739908,0.0005661225,0.00042513397,0.000381754,0.0008146305,0.0009590504,0.00094790396,0.0024613475],"category_scores_gemma":[0.002929891,0.00033664855,0.0004794375,0.00027913426,0.0010054125,0.0012488348,0.0011037256,0.0012907604,0.00040019408],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003352984,0.000031598018,0.0005082949,0.000051106133,0.000048705995,0.00014936349,0.00014541952,0.8283646,0.0072515905,0.13573086,0.0015216939,0.026163248],"study_design_scores_gemma":[0.0000035950454,0.000007327845,0.00004421457,0.0000026666473,0.0000035565727,0.00001120715,0.0000066257585,0.97324175,0.00035163644,0.025713868,0.00060971285,0.000003774898],"about_ca_topic_score_codex":0.0021737348,"about_ca_topic_score_gemma":0.0023470782,"teacher_disagreement_score":0.0024613475,"about_ca_system_score_codex":0.0011755466,"about_ca_system_score_gemma":0.0005150496,"threshold_uncertainty_score":0.008529246},"labels":[],"label_agreement":null},{"id":"W2962770186","doi":"10.18653/v1/d15-1179","title":"Fast, Flexible Models for Discovering Topic Correlation across Weakly-Related Collections","year":2015,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; John Templeton Foundation; National Science Foundation","keywords":"Computer science; Correlation; Information retrieval; Mathematics","score_opus":0.06405536907975934,"score_gpt":0.3114599321620022,"score_spread":0.24740456308224285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2962770186","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03139471,0.0019041036,0.9609155,0.0007505859,0.0001231149,0.00024477017,0.0015929687,0.002072631,0.0010016578],"genre_scores_gemma":[0.49386936,0.003297227,0.47899148,0.0006290411,0.0012973805,0.0013153059,0.013365319,0.0009952975,0.0062396075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99407023,0.0022653495,0.0004439308,0.001923162,0.00086104247,0.00043623627],"domain_scores_gemma":[0.9730987,0.018208731,0.0015320021,0.0044083786,0.0020591354,0.0006929367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012212998,0.0020405545,0.0032614851,0.004855583,0.0020131252,0.0052188803,0.005731986,0.0029858495,0.0019850668],"category_scores_gemma":[0.03701021,0.0021812962,0.0030108627,0.0069892155,0.0015265028,0.0113176275,0.0043846364,0.006540492,0.003154216],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016002301,0.00089959527,0.03199399,0.0012184293,0.0015350434,0.00051117974,0.0021864658,0.28212523,0.009523093,0.051022075,0.03802875,0.5793559],"study_design_scores_gemma":[0.00005842756,0.00006744612,0.0019759494,0.00004757013,0.000116721305,0.0001473548,0.00016251067,0.9404824,0.0011682225,0.052112896,0.0036138098,0.000046664816],"about_ca_topic_score_codex":0.009753645,"about_ca_topic_score_gemma":0.015885707,"teacher_disagreement_score":0.012212998,"about_ca_system_score_codex":0.0017249368,"about_ca_system_score_gemma":0.0028192876,"threshold_uncertainty_score":0.06458926},"labels":[],"label_agreement":null},{"id":"W2962810136","doi":"10.18653/v1/d17-1043","title":"Piecewise Latent Variables for Neural Variational Text Processing","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Nuance Foundation; Canadian Institute for Advanced Research","keywords":"Latent variable; Piecewise; Computer science; Artificial intelligence; Prior probability; Latent variable model; Inference; Probabilistic latent semantic analysis; Gaussian; Natural language; Exponential family; Artificial neural network; Algorithm; Machine learning; Mathematics; Bayesian probability","score_opus":0.051327734419385346,"score_gpt":0.2843952514319339,"score_spread":0.23306751701254855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2962810136","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002109811,0.0003983141,0.99556893,0.00034843868,0.000036198544,0.00001863823,0.00012421416,0.00024095718,0.0011545033],"genre_scores_gemma":[0.33556005,0.002110805,0.64124185,0.0006106393,0.00045611264,0.00060161954,0.0013824546,0.0009858479,0.017050678],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99933136,0.00030185017,0.00003099656,0.00014664036,0.00013711142,0.000052104046],"domain_scores_gemma":[0.9975987,0.0018103627,0.00015517381,0.0001910499,0.00016211889,0.00008260261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016165513,0.0008537999,0.00076148077,0.0009352811,0.0005700391,0.0012961563,0.0019404163,0.0014160696,0.006582571],"category_scores_gemma":[0.008021561,0.0007436207,0.0010329477,0.0013912176,0.0013986137,0.0031265363,0.0019293728,0.0035277156,0.0012877494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058006943,0.000030278883,0.00036214394,0.00014216195,0.000052404102,0.000064730186,0.00015252217,0.39071313,0.0017390876,0.5589775,0.0042529423,0.043455143],"study_design_scores_gemma":[0.00000381217,0.0000043937434,0.000041397787,0.0000074381715,0.0000035171977,0.00000859048,0.0000054048196,0.83443254,0.00020134123,0.16385445,0.0014313633,0.0000056843],"about_ca_topic_score_codex":0.0057353945,"about_ca_topic_score_gemma":0.0064819804,"teacher_disagreement_score":0.006582571,"about_ca_system_score_codex":0.0019211309,"about_ca_system_score_gemma":0.001173565,"threshold_uncertainty_score":0.022020876},"labels":[],"label_agreement":null},{"id":"W2962852556","doi":"","title":"Testing APSyn against Vector Cosine on Similarity Estimation","year":2016,"lang":"en","type":"preprint","venue":"Waseda University Repository (Waseda University)","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited; University of Oxford; University of Wisconsin-Madison","keywords":"Cosine similarity; Similarity (geometry); Relevance (law); Weighting; Context (archaeology); Word (group theory); Metric (unit); Computer science; Measure (data warehouse); Intersection (aeronautics); Similarity measure; Vector space model; Artificial intelligence; Trigonometric functions; Semantic similarity; Estimation; Pattern recognition (psychology); Task (project management); Mathematics; Natural language processing; Data mining; Image (mathematics)","score_opus":0.024139572270494652,"score_gpt":0.19713452692286537,"score_spread":0.17299495465237072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2962852556","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7901969,0.012124401,0.16078813,0.0018632547,0.0020308606,0.0007237211,0.0062947967,0.005882504,0.020095497],"genre_scores_gemma":[0.9246397,0.00071106915,0.06364013,0.00034368256,0.00035496504,0.0003196784,0.007909831,0.00037336536,0.0017075444],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.959202,0.02341653,0.0030317537,0.005664831,0.0076890057,0.0009959368],"domain_scores_gemma":[0.89632857,0.08080933,0.0032678307,0.011546222,0.006239257,0.0018088305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030040568,0.0024358078,0.002338748,0.005737766,0.0013441517,0.0033451368,0.003367864,0.0034201154,0.004408792],"category_scores_gemma":[0.1280428,0.00047464794,0.001387514,0.005314289,0.0022392822,0.010094165,0.0059234635,0.0023590063,0.0027236135],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009468916,0.0021073967,0.13928321,0.0026709975,0.0038091268,0.00045274355,0.00079802796,0.12623101,0.0069538425,0.018414007,0.03532473,0.65448594],"study_design_scores_gemma":[0.0005320146,0.004499413,0.02729315,0.00022937584,0.0004424254,0.0010162706,0.0013661956,0.92873186,0.0083024455,0.019578595,0.007855703,0.00015243512],"about_ca_topic_score_codex":0.0040783957,"about_ca_topic_score_gemma":0.0030508183,"teacher_disagreement_score":0.030040568,"about_ca_system_score_codex":0.0011369425,"about_ca_system_score_gemma":0.0017216754,"threshold_uncertainty_score":0.15887159},"labels":[],"label_agreement":null},{"id":"W2962876041","doi":"","title":"An Exploration of Softmax Alternatives Belonging to the Spherical Loss Family","year":2016,"lang":"en","type":"article","venue":"International Conference on Learning Representations","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Softmax function; Probabilistic logic; Function (biology); Categorical variable; MNIST database; Computer science; Range (aeronautics); Mathematics; Algorithm; Artificial neural network; Artificial intelligence; Pattern recognition (psychology); Discrete mathematics; Machine learning; Engineering","score_opus":0.10429671052686768,"score_gpt":0.37495094648723215,"score_spread":0.2706542359603645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2962876041","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026206324,0.0013991396,0.9643712,0.0014856552,0.00006074282,0.000058397974,0.00009479252,0.00059051684,0.0057331547],"genre_scores_gemma":[0.7324584,0.002283612,0.2556267,0.0013948715,0.00022007209,0.0003368707,0.0005112169,0.00068506366,0.006483323],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9969199,0.0018819305,0.00016631823,0.0003123545,0.0005415376,0.00017798186],"domain_scores_gemma":[0.99450064,0.003863962,0.00030847624,0.0005327034,0.0005707045,0.00022352862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007464447,0.0016352456,0.0011225989,0.0012305962,0.00053175085,0.0024032714,0.002021552,0.0017757685,0.0031889235],"category_scores_gemma":[0.017089667,0.00051039923,0.0012908378,0.0012411582,0.001624844,0.004285737,0.003324156,0.003243951,0.001385185],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074717664,0.00033011258,0.0019182204,0.00038532726,0.00018108207,0.00025051748,0.00029626497,0.5057165,0.0057495707,0.22384587,0.008905797,0.25167367],"study_design_scores_gemma":[0.000028625489,0.000126225,0.00027628534,0.00006327845,0.000021171601,0.000107273794,0.000055542678,0.9342073,0.0017338791,0.061275203,0.0020837118,0.000021560762],"about_ca_topic_score_codex":0.0015684456,"about_ca_topic_score_gemma":0.0016462784,"teacher_disagreement_score":0.007464447,"about_ca_system_score_codex":0.0015902641,"about_ca_system_score_gemma":0.001124616,"threshold_uncertainty_score":0.039476275},"labels":[],"label_agreement":null},{"id":"W2962883855","doi":"10.1609/aaai.v30i1.9883","title":"Building End-To-End Dialogue Systems Using Generative Hierarchical Neural Network Models","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1725,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Canadian Institute for Advanced Research","keywords":"Generative grammar; Computer science; Bootstrapping (finance); Word (group theory); Artificial intelligence; End-to-end principle; Artificial neural network; Language model; Domain (mathematical analysis); Encoder; Recurrent neural network; Task (project management); Natural language processing; Generative model; Linguistics","score_opus":0.06428255094117762,"score_gpt":0.2761510849675109,"score_spread":0.21186853402633327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2962883855","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030636262,0.00015955998,0.9621503,0.00016597024,0.000059016536,0.00009642297,0.00013073652,0.0049470593,0.0016546871],"genre_scores_gemma":[0.6042902,0.00016691444,0.38842377,0.00027431865,0.000056452835,0.00037040774,0.00090968475,0.000541444,0.004966798],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998995,0.00040780194,0.00004285455,0.0003666259,0.00011518181,0.00007258062],"domain_scores_gemma":[0.9977677,0.0015198324,0.000095767595,0.00022744248,0.00026865065,0.00012055108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017372164,0.0009312491,0.00087408663,0.00042000052,0.0005862525,0.0012218113,0.0016342206,0.0013991912,0.0029556227],"category_scores_gemma":[0.0060550794,0.000776016,0.0010385296,0.00031947333,0.0007866192,0.0026551005,0.0020538638,0.0020586632,0.0021635816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004030473,0.0003926896,0.0015497041,0.00034076883,0.00024025956,0.00046955294,0.0013250438,0.7318091,0.036467943,0.020855565,0.004662604,0.2014837],"study_design_scores_gemma":[0.000008770001,0.000021833725,0.000059879007,0.000005056487,0.000010568294,0.000018558181,0.000024928953,0.99202216,0.0024827977,0.0047798133,0.00055756234,0.00000810927],"about_ca_topic_score_codex":0.0030292892,"about_ca_topic_score_gemma":0.0050488194,"teacher_disagreement_score":0.0030292892,"about_ca_system_score_codex":0.00078091235,"about_ca_system_score_gemma":0.00082280964,"threshold_uncertainty_score":0.0098875165},"labels":[],"label_agreement":null},{"id":"W2963085936","doi":"10.18653/v1/n18-1002","title":"Neural Fine-Grained Entity Type Classification with Hierarchy-Aware Loss","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Heuristics; Artificial intelligence; Normalization (sociology); Sentence; Hierarchy; Artificial neural network; Machine learning; Task (project management); Natural language processing","score_opus":0.03957987951996678,"score_gpt":0.26872099550189044,"score_spread":0.22914111598192366,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963085936","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.115358025,0.0024181905,0.85695285,0.0011095944,0.00041828415,0.00020833455,0.0025760378,0.016185058,0.0047735237],"genre_scores_gemma":[0.65731966,0.000781435,0.30379906,0.00086633844,0.00045124505,0.00036667057,0.014567807,0.00060676166,0.02124106],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918884,0.00015544872,0.00005474747,0.00030488576,0.0001584299,0.00013766203],"domain_scores_gemma":[0.9984353,0.00062188174,0.00013717751,0.0003774442,0.00034666102,0.000081361926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018404506,0.0015609324,0.0015058087,0.0018264959,0.0006261348,0.0011898436,0.002920856,0.0022790995,0.0030379016],"category_scores_gemma":[0.00369695,0.0005029416,0.0010893169,0.0024518194,0.00069110934,0.0043192953,0.0017001939,0.0028184392,0.0024423422],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006530934,0.000505271,0.0037311711,0.00024354075,0.00013815923,0.00016958034,0.00019307245,0.17907228,0.008665637,0.0071729887,0.032584213,0.76687104],"study_design_scores_gemma":[0.00002460579,0.000055436733,0.0004710031,0.0000135601795,0.000024179213,0.00004076561,0.000030013904,0.98658216,0.0028426964,0.008458488,0.0014448799,0.000012187768],"about_ca_topic_score_codex":0.007495348,"about_ca_topic_score_gemma":0.011224221,"teacher_disagreement_score":0.007495348,"about_ca_system_score_codex":0.0014249445,"about_ca_system_score_gemma":0.0010814561,"threshold_uncertainty_score":0.014903426},"labels":[],"label_agreement":null},{"id":"W2963172229","doi":"10.1609/aaai.v33i01.33016762","title":"Contextualized Non-Local Neural Networks for Sequence Learning","year":2019,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Science and Technology Commission of Shanghai Municipality; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Interpretability; Computer science; Artificial intelligence; Leverage (statistics); Artificial neural network; Sentence; Machine learning; Sequence learning; Transformer; Natural language processing","score_opus":0.09166343411648817,"score_gpt":0.31093158964930734,"score_spread":0.21926815553281917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963172229","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077188383,0.0010527109,0.9887905,0.00023687114,0.000043171596,0.000030010542,0.00015843255,0.0009475252,0.0010219107],"genre_scores_gemma":[0.5728072,0.0020326104,0.41682658,0.00044810557,0.0002613201,0.00035845026,0.0013106134,0.00027951007,0.0056757196],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99953806,0.00016540938,0.000023908831,0.00016632843,0.000071786955,0.000034565182],"domain_scores_gemma":[0.9990741,0.0005275928,0.00008632776,0.00014855638,0.00012720523,0.000036221416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00089664455,0.00083108386,0.0007821267,0.0008723526,0.00031716633,0.00059026893,0.0014819814,0.0010099361,0.0032473],"category_scores_gemma":[0.0033066815,0.00040847823,0.0006150447,0.0012578795,0.0006447702,0.0018747939,0.0009053024,0.0016387617,0.00082988176],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014917142,0.00008649911,0.0008313332,0.00022373644,0.00010671369,0.00013661469,0.00012962673,0.7261063,0.0058028414,0.05202677,0.0036092703,0.21079113],"study_design_scores_gemma":[0.0000036420533,0.0000177869,0.000087102475,0.0000063098146,0.0000066325742,0.000010140385,0.000004719815,0.9747354,0.0005168628,0.023879444,0.0007271992,0.000004782538],"about_ca_topic_score_codex":0.0061133546,"about_ca_topic_score_gemma":0.010377844,"teacher_disagreement_score":0.0061133546,"about_ca_system_score_codex":0.0010575046,"about_ca_system_score_gemma":0.0007011627,"threshold_uncertainty_score":0.012155533},"labels":[],"label_agreement":null},{"id":"W2963172394","doi":"10.1109/icassp.2019.8682634","title":"Why Do Neural Dialog Systems Generate Short and Meaningless Replies? a Comparison between Dialog and Translation","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Natural Science Foundation of China","keywords":"Dialog box; Computer science; Utterance; Natural language processing; Machine translation; Randomness; Dialog system; Sequence (biology); Artificial intelligence; Translation (biology); Speech recognition; Conjecture; World Wide Web; Mathematics","score_opus":0.05761809124299732,"score_gpt":0.26601110357707974,"score_spread":0.20839301233408242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963172394","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58471614,0.0035181365,0.38236335,0.0037305343,0.00036613055,0.00030612163,0.0013371903,0.0035386642,0.020123731],"genre_scores_gemma":[0.9666745,0.00032209727,0.029810274,0.0004036989,0.00008575751,0.00015111807,0.0008651398,0.0001653341,0.00152218],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954608,0.0027717121,0.00024498478,0.00085383985,0.00048966287,0.00017910569],"domain_scores_gemma":[0.9696474,0.02384852,0.0010532083,0.0029299904,0.0020226815,0.0004980579],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007172694,0.0008076675,0.00090687204,0.00088843703,0.00090777065,0.001630367,0.0009923682,0.0017786667,0.0038806868],"category_scores_gemma":[0.045900553,0.00043807604,0.0006126781,0.00070110156,0.001325341,0.003955623,0.001326355,0.001295177,0.001765287],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0054420796,0.0008905565,0.034837782,0.0032111993,0.00090945326,0.001050729,0.0065372875,0.24263622,0.06897146,0.05538908,0.011678321,0.5684458],"study_design_scores_gemma":[0.00038120503,0.0011824851,0.021149848,0.00017716306,0.00027250592,0.0009091556,0.0015351682,0.8134978,0.041584417,0.108748496,0.010341594,0.00022016573],"about_ca_topic_score_codex":0.0014724841,"about_ca_topic_score_gemma":0.0013781033,"teacher_disagreement_score":0.007172694,"about_ca_system_score_codex":0.0008236902,"about_ca_system_score_gemma":0.0008139446,"threshold_uncertainty_score":0.03793323},"labels":[],"label_agreement":null},{"id":"W2963185998","doi":"","title":"Plan, Attend, Generate: Planning for Sequence-to-Sequence Models","year":2017,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; Université de Montréal","funders":"","keywords":"Computer science; Sequence (biology); Plan (archaeology); Task (project management); Baseline (sea); Artificial intelligence; Character (mathematics); Mechanism (biology); Machine learning; Mathematics; Engineering","score_opus":0.19433423543765413,"score_gpt":0.3457445669782557,"score_spread":0.15141033154060154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963185998","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028524017,0.00032060174,0.9635057,0.0010307182,0.00007868178,0.0001062767,0.00034547856,0.0018195183,0.0042689415],"genre_scores_gemma":[0.7408161,0.00043090532,0.2499529,0.0003799875,0.000079288984,0.00045771943,0.0006064133,0.00029904125,0.006977751],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99961203,0.0001559005,0.000017824486,0.000111533474,0.000054962897,0.000047792775],"domain_scores_gemma":[0.99760324,0.0018072623,0.00012293358,0.00022605143,0.00012507543,0.00011539997],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013023063,0.000790839,0.00074914546,0.0004320623,0.00044994312,0.0010533645,0.0018653168,0.0014741963,0.0069642933],"category_scores_gemma":[0.0061695883,0.0006023369,0.00078017113,0.0005725307,0.0009787976,0.0023838254,0.0010336398,0.0023030338,0.00088948524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011369772,0.0000545512,0.0005954764,0.00007574782,0.000033317046,0.00008992429,0.00012979277,0.91549885,0.0009076377,0.043008022,0.0026403414,0.036852654],"study_design_scores_gemma":[0.000008497728,0.0000116122255,0.000028669741,0.0000029685154,0.0000045603283,0.000007650599,0.0000047959675,0.98217726,0.00025542802,0.017114729,0.00038102266,0.0000029215428],"about_ca_topic_score_codex":0.01110304,"about_ca_topic_score_gemma":0.014931319,"teacher_disagreement_score":0.01110304,"about_ca_system_score_codex":0.0015575961,"about_ca_system_score_gemma":0.0017234833,"threshold_uncertainty_score":0.023297906},"labels":[],"label_agreement":null},{"id":"W2963241005","doi":"10.1145/3322640.3326711","title":"A Reliable and Accurate Multiple Choice Question Answering System for Due Diligence","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cisco Systems (Canada)","funders":"","keywords":"Question answering; Computer science; Due diligence; Scarcity; Artificial intelligence; Classifier (UML); Task (project management); Machine learning; Oversampling; Bandwidth (computing); Finance","score_opus":0.021179623793798882,"score_gpt":0.2575918210204401,"score_spread":0.2364121972266412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963241005","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.093364,0.00093594746,0.81233805,0.0025251915,0.00043845773,0.001047468,0.005134393,0.07796198,0.0062545743],"genre_scores_gemma":[0.45058393,0.00028490985,0.5254179,0.0008567972,0.0004183792,0.0008233745,0.014085777,0.00054992025,0.006979011],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968353,0.0007772231,0.00028898072,0.00097412284,0.0009096365,0.00021467768],"domain_scores_gemma":[0.99250257,0.003409779,0.0005461411,0.0008961972,0.002268547,0.00037666145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039840634,0.0009424471,0.0015976181,0.002103213,0.0010234762,0.0017303283,0.002099779,0.0031638322,0.0057826503],"category_scores_gemma":[0.01359957,0.00035060954,0.00072532636,0.0012846342,0.00040786772,0.0037587439,0.0017082333,0.0017801403,0.005941562],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010010431,0.0012720969,0.014805615,0.00063285884,0.00016102563,0.0008491563,0.0011657747,0.01856472,0.07764192,0.010370869,0.102191634,0.7713433],"study_design_scores_gemma":[0.00011406176,0.00027645982,0.0069463393,0.000054494263,0.000072220784,0.0005364918,0.00032165294,0.90962094,0.03759063,0.012388474,0.031962287,0.00011582483],"about_ca_topic_score_codex":0.003613643,"about_ca_topic_score_gemma":0.003108896,"teacher_disagreement_score":0.0057826503,"about_ca_system_score_codex":0.0010037629,"about_ca_system_score_gemma":0.0014641316,"threshold_uncertainty_score":0.021070004},"labels":[],"label_agreement":null},{"id":"W2963301888","doi":"10.18653/v1/n18-4017","title":"Training a Ranking Function for Open-Domain Question Answering","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Advanced Research Projects Agency; Tencent; Defense Advanced Research Projects Agency; Canadian Institute for Advanced Research; Samsung; Nvidia","keywords":"Computer science; Question answering; Artificial intelligence; Natural language processing; Paragraph; Reading (process); Reading comprehension; Relevance (law); Ranking (information retrieval); Similarity (geometry); Task (project management); Domain (mathematical analysis); Information retrieval; Linguistics; World Wide Web","score_opus":0.07747956067232864,"score_gpt":0.31364003545740976,"score_spread":0.23616047478508112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963301888","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08384327,0.0023024345,0.90151876,0.00082497194,0.0001871608,0.00021185819,0.0005394733,0.0066320323,0.0039399723],"genre_scores_gemma":[0.67451245,0.000576982,0.3119742,0.00056163146,0.00037718468,0.0004239853,0.0033404413,0.00043346113,0.0077995555],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985757,0.0006301322,0.00009146707,0.00032294472,0.00021947939,0.00016029578],"domain_scores_gemma":[0.99594814,0.0027729159,0.000163446,0.0002948553,0.0006858507,0.00013479339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002836384,0.0012981737,0.0014353216,0.0014951605,0.00053590926,0.00089218444,0.0018314482,0.0025370617,0.0058613443],"category_scores_gemma":[0.0095644705,0.000434116,0.00079276005,0.0010256794,0.00048049,0.0019607926,0.0009567281,0.001819315,0.0029913012],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044288323,0.00062873686,0.0030973083,0.0003427968,0.0001126571,0.00015032332,0.00018075496,0.18477182,0.0072889663,0.0073467293,0.018846422,0.77679056],"study_design_scores_gemma":[0.000026006099,0.00009342275,0.00039237956,0.000010987652,0.000013169625,0.00003193699,0.000024075758,0.99292505,0.0013737213,0.0041619106,0.0009385263,0.000008899889],"about_ca_topic_score_codex":0.0038192843,"about_ca_topic_score_gemma":0.005213969,"teacher_disagreement_score":0.0058613443,"about_ca_system_score_codex":0.00094399066,"about_ca_system_score_gemma":0.0010036391,"threshold_uncertainty_score":0.01960814},"labels":[],"label_agreement":null},{"id":"W2963386218","doi":"","title":"A Structured Self-Attentive Sentence Embedding.","year":2017,"lang":"en","type":"article","venue":"International Conference on Learning Representations","topic":"Topic Modeling","field":"Computer Science","cited_by":365,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Embedding; Sentence; Computer science; Natural language processing; Artificial intelligence; Logical consequence; Regularization (linguistics)","score_opus":0.05458417324847135,"score_gpt":0.3661503131841141,"score_spread":0.31156613993564275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963386218","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0269405,0.0008884918,0.9651569,0.00058559567,0.00038660684,0.00013061582,0.0009334377,0.0017715411,0.0032063294],"genre_scores_gemma":[0.5320228,0.0010985476,0.44544572,0.00061754714,0.00040204608,0.00029680436,0.003976705,0.00034639108,0.01579355],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99974674,0.00008131277,0.000015302248,0.00009648712,0.000043541444,0.000016654007],"domain_scores_gemma":[0.9995179,0.00019600929,0.00006180477,0.00008805542,0.00010644845,0.000029737033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00038450974,0.00067327364,0.00028888558,0.0006852043,0.00015726266,0.00054339715,0.0007281756,0.00067735463,0.0031482023],"category_scores_gemma":[0.002068233,0.00025170078,0.0006085125,0.0004781448,0.00026133755,0.0018561205,0.00064956327,0.00090409606,0.0013714952],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035737306,0.00037777558,0.0033694443,0.0006048552,0.00022358846,0.00034373344,0.0007884331,0.063860774,0.07697031,0.039438844,0.029190587,0.7844743],"study_design_scores_gemma":[0.000015049276,0.00017462988,0.0015821961,0.000054336742,0.00006841644,0.00023391919,0.00007608756,0.9481597,0.012092609,0.023553086,0.013960749,0.000029140752],"about_ca_topic_score_codex":0.0011365489,"about_ca_topic_score_gemma":0.0021110123,"teacher_disagreement_score":0.0031482023,"about_ca_system_score_codex":0.0003150561,"about_ca_system_score_gemma":0.00034844043,"threshold_uncertainty_score":0.010531783},"labels":[],"label_agreement":null},{"id":"W2963412005","doi":"10.18653/v1/d16-1233","title":"Conditional Generation and Snapshot Learning in Neural Dialogue Systems","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Institute for Catastrophic Loss Reduction","keywords":"Snapshot (computer storage); Computer science; Artificial intelligence; Artificial neural network; Natural language processing; Cognitive science; Machine learning; Psychology; Operating system","score_opus":0.038899565773714104,"score_gpt":0.23756521787474763,"score_spread":0.19866565210103354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963412005","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.065294564,0.0018013513,0.9274912,0.0010256353,0.00012740154,0.00006465286,0.00025474,0.0014016047,0.0025387867],"genre_scores_gemma":[0.88706124,0.0003630525,0.1077898,0.00021393396,0.00013997081,0.00015942768,0.00063572725,0.00028405126,0.0033528088],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983359,0.0010105566,0.000069561065,0.00033290903,0.0001261805,0.00012503915],"domain_scores_gemma":[0.98266023,0.015074691,0.00035031358,0.0008615293,0.000655317,0.00039793234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003974202,0.0006634598,0.0012879916,0.0008274071,0.0007012218,0.0015101604,0.0019880459,0.0013382917,0.003994774],"category_scores_gemma":[0.021163976,0.0008157483,0.0005354696,0.000690865,0.0012554205,0.0042767646,0.002944919,0.0023386078,0.0005683884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008491372,0.00017427391,0.0032278998,0.00027969037,0.00016126788,0.00021632592,0.0007721269,0.6561083,0.0015359868,0.13948667,0.007296574,0.18989177],"study_design_scores_gemma":[0.000016128351,0.000017931903,0.0001324248,0.0000070198103,0.0000076255733,0.000016100643,0.000015478478,0.93074816,0.0002796853,0.06847088,0.00027991613,0.000008608237],"about_ca_topic_score_codex":0.0045227674,"about_ca_topic_score_gemma":0.0066433214,"teacher_disagreement_score":0.0045227674,"about_ca_system_score_codex":0.0011634645,"about_ca_system_score_gemma":0.000899228,"threshold_uncertainty_score":0.02101785},"labels":[],"label_agreement":null},{"id":"W2963494066","doi":"10.18653/v1/p17-2059","title":"A Deep Network with Visual Text Composition Behavior","year":2017,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Principle of compositionality; Computer science; Sentence; Artificial intelligence; Contrast (vision); Natural language processing; Composition (language); Layer (electronics); Linguistics","score_opus":0.027882928834677967,"score_gpt":0.2943428115788486,"score_spread":0.26645988274417065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963494066","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.095133096,0.0005317459,0.8870098,0.0012589581,0.00023066223,0.00008262266,0.0008919895,0.006734628,0.008126457],"genre_scores_gemma":[0.7245907,0.0003596511,0.2539304,0.00060146657,0.000116544004,0.00017573858,0.0019480381,0.0003716935,0.017905746],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998605,0.000017245657,0.0000060102993,0.000066530265,0.000024140381,0.00002549516],"domain_scores_gemma":[0.9997843,0.00006693229,0.000021447242,0.000041500465,0.00005327282,0.000032533666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00026946195,0.0007579991,0.00031203814,0.00044521628,0.00034216276,0.00069818547,0.0010015272,0.00090452563,0.0038166873],"category_scores_gemma":[0.0013108355,0.00042709737,0.0005079381,0.0004201945,0.0004807395,0.0017919154,0.0010016633,0.0011284803,0.0012358745],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007545799,0.0004151104,0.0033665735,0.00031604306,0.000158504,0.00046776517,0.00037088242,0.31818244,0.15490544,0.031724617,0.018962558,0.4703755],"study_design_scores_gemma":[0.000015597445,0.00003124975,0.00020403611,0.00000926375,0.000018246637,0.0000275711,0.000012409812,0.98249567,0.007258188,0.008208494,0.0017123974,0.0000068721984],"about_ca_topic_score_codex":0.00542559,"about_ca_topic_score_gemma":0.0096345255,"teacher_disagreement_score":0.00542559,"about_ca_system_score_codex":0.0007730235,"about_ca_system_score_gemma":0.0006535411,"threshold_uncertainty_score":0.01276809},"labels":[],"label_agreement":null},{"id":"W2963494503","doi":"10.1609/aaai.v33i01.33013526","title":"Improved Knowledge Graph Embedding Using Background Taxonomic Information","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Embedding; Knowledge graph; Computer science; Theoretical computer science; Graph; Mathematics; Artificial intelligence","score_opus":0.03363265289143851,"score_gpt":0.2717997519343149,"score_spread":0.23816709904287636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963494503","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030567115,0.00036449972,0.9635743,0.0005035035,0.00004046696,0.00007053788,0.00089664874,0.0021815593,0.0018013698],"genre_scores_gemma":[0.43379086,0.00072849554,0.54929435,0.0002921506,0.00011384366,0.00017472924,0.007373908,0.00067518844,0.007556533],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99877566,0.00033110988,0.00007499018,0.0004463762,0.00027832246,0.000093496135],"domain_scores_gemma":[0.9954514,0.0019279373,0.0003080635,0.0017528415,0.0004068537,0.00015288212],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012207074,0.0010065794,0.0010016279,0.0017259627,0.00060840027,0.0018583296,0.0015733608,0.0012890581,0.0034006743],"category_scores_gemma":[0.008158849,0.0005517201,0.001571419,0.0022565178,0.00076324394,0.0077970047,0.0026801303,0.0023452847,0.0012417727],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030272338,0.000489266,0.0030822463,0.0005961262,0.00022494658,0.00044969865,0.0009802454,0.29021013,0.0151349185,0.11560618,0.021258408,0.5516651],"study_design_scores_gemma":[0.000014473127,0.000028955601,0.00030049525,0.000027670712,0.00004253483,0.000099541394,0.00009201456,0.91835225,0.0022220158,0.0749045,0.0038998318,0.000015673797],"about_ca_topic_score_codex":0.004811972,"about_ca_topic_score_gemma":0.0076655988,"teacher_disagreement_score":0.004811972,"about_ca_system_score_codex":0.0008771682,"about_ca_system_score_gemma":0.0010962308,"threshold_uncertainty_score":0.0113764405},"labels":[],"label_agreement":null},{"id":"W2963506530","doi":"10.18653/v1/p19-1602","title":"Generating Sentences from Disentangled Syntactic and Semantic Spaces","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Natural Science Foundation of China; National Science Foundation","keywords":"Computer science; Paraphrase; Natural language processing; Artificial intelligence; Syntax; Language model; Latent semantic analysis","score_opus":0.027091851340920344,"score_gpt":0.25425543778985954,"score_spread":0.2271635864489392,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963506530","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044227425,0.0003640893,0.9508805,0.00034102658,0.00010094383,0.00013603517,0.00052814087,0.001597644,0.0018241769],"genre_scores_gemma":[0.6223118,0.0004296664,0.36755183,0.00027923554,0.0001581581,0.00046094562,0.0032835219,0.0005241012,0.0050007817],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99940205,0.00026641157,0.000033788616,0.0001642133,0.000092495786,0.000041090538],"domain_scores_gemma":[0.99850464,0.0010226377,0.00008269166,0.00015423603,0.0001853648,0.000050421775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007793315,0.0009610806,0.0006598045,0.00061535806,0.00032005622,0.00055401766,0.0008844803,0.00089656084,0.002356681],"category_scores_gemma":[0.0034025947,0.00041241912,0.0012717921,0.0005772053,0.0004582509,0.0014380867,0.0010885042,0.001295586,0.000850795],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044360803,0.00034170027,0.0025975737,0.00077924563,0.0002506101,0.001005395,0.0008567821,0.27900043,0.08244657,0.080136925,0.016871367,0.53526986],"study_design_scores_gemma":[0.000033078202,0.00005074746,0.00026228486,0.000014687523,0.00003370167,0.00010653993,0.00004022315,0.9623027,0.010233937,0.02506603,0.0018412163,0.000014930757],"about_ca_topic_score_codex":0.0009826968,"about_ca_topic_score_gemma":0.0021500008,"teacher_disagreement_score":0.002356681,"about_ca_system_score_codex":0.00037271727,"about_ca_system_score_gemma":0.00089087035,"threshold_uncertainty_score":0.007883847},"labels":[],"label_agreement":null},{"id":"W2963539636","doi":"","title":"ClaC: Semantic Relatedness of Words and Phrases","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Semantic similarity; Computer science; Natural language processing; Similarity (geometry); Artificial intelligence; Semantics (computer science); Task (project management); Metric (unit); Distributional semantics; Semantic computing; Semantic compression; Semantic technology; Semantic Web; Programming language","score_opus":0.014849971191806121,"score_gpt":0.2227922787747524,"score_spread":0.2079423075829463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963539636","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35654804,0.007401253,0.54788065,0.0012403113,0.00059398613,0.0020991415,0.037626393,0.013957613,0.032652684],"genre_scores_gemma":[0.7770754,0.0009576063,0.1848194,0.00025834126,0.0003177967,0.001220987,0.03066654,0.0009815561,0.0037023379],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912689,0.002391132,0.00084944355,0.0016416122,0.0034662697,0.00038278292],"domain_scores_gemma":[0.9823315,0.007853151,0.002493534,0.0024026027,0.0042863693,0.0006328285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048265117,0.0016840309,0.001385829,0.019104818,0.0016846892,0.0027211173,0.0020005524,0.0021004265,0.005378756],"category_scores_gemma":[0.026876839,0.00036131009,0.0011597571,0.011337113,0.0013041843,0.0056356657,0.0029692464,0.0014365044,0.0034356571],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015737144,0.0008818445,0.08847655,0.003153886,0.000995384,0.0006515629,0.002175564,0.02193932,0.09361881,0.033309754,0.044159956,0.7090637],"study_design_scores_gemma":[0.00020482685,0.0017262294,0.18493864,0.0004185036,0.0006395632,0.004875826,0.0019049378,0.5702154,0.066963,0.09430075,0.07321347,0.0005987993],"about_ca_topic_score_codex":0.0060075494,"about_ca_topic_score_gemma":0.0068835267,"teacher_disagreement_score":0.019104818,"about_ca_system_score_codex":0.0016220293,"about_ca_system_score_gemma":0.0017277974,"threshold_uncertainty_score":0.025525331},"labels":[],"label_agreement":null},{"id":"W2963595285","doi":"","title":"Simple Search Algorithms on Semantic Networks Learned from Language Use","year":2016,"lang":"en","type":"article","venue":"eScholarship (California Digital Library)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Semantic memory; Artificial intelligence; Simple (philosophy); Task (project management); Natural language processing; Structuring; Theoretical computer science; Machine learning; Cognition; Psychology","score_opus":0.03446120323187399,"score_gpt":0.2433639120874294,"score_spread":0.20890270885555542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963595285","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3522827,0.00089255464,0.6379988,0.0011154269,0.000060131144,0.00015864338,0.000646913,0.0009859772,0.005858887],"genre_scores_gemma":[0.8899744,0.00042958616,0.10519108,0.00016227718,0.00006025624,0.00020072908,0.0008386884,0.00011700208,0.0030258445],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995704,0.00017375535,0.000030278698,0.00013935933,0.000044904795,0.000041310916],"domain_scores_gemma":[0.9948597,0.0041745817,0.0003618787,0.00033267177,0.00018050209,0.00009069528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014712582,0.00075979397,0.000784614,0.0017373759,0.00035312187,0.0009337563,0.001278282,0.0012358516,0.0031185006],"category_scores_gemma":[0.013020778,0.00043605457,0.0006819548,0.0012062076,0.0009929962,0.004429629,0.00079540466,0.0010585402,0.00040986194],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032040413,0.00018522941,0.005708355,0.00030276002,0.00015537249,0.00014689063,0.00033999447,0.779357,0.0024064675,0.075492166,0.0029947045,0.13259056],"study_design_scores_gemma":[0.00002305472,0.000025037667,0.0003511523,0.00001511318,0.0000099424415,0.000027264603,0.000021158055,0.9341296,0.00029950464,0.06483204,0.0002589931,0.000007123944],"about_ca_topic_score_codex":0.0036384042,"about_ca_topic_score_gemma":0.006916007,"teacher_disagreement_score":0.0036384042,"about_ca_system_score_codex":0.0011306421,"about_ca_system_score_gemma":0.0006709708,"threshold_uncertainty_score":0.010432482},"labels":[],"label_agreement":null},{"id":"W2963603213","doi":"","title":"Efficient Exact Gradient Update for training Deep Networks with Very Large Sparse Targets.","year":2015,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal; Canadian Institute for Advanced Research","funders":"","keywords":"Softmax function; Backpropagation; Computer science; Artificial neural network; Speedup; Dimension (graph theory); Algorithm; Deep learning; Computation; Sparse matrix; Gradient descent; Artificial intelligence; Mathematics; Parallel computing","score_opus":0.04118044383817406,"score_gpt":0.2430355192364063,"score_spread":0.20185507539823222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963603213","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007815802,0.00015021938,0.9876605,0.00018494207,0.000045735353,0.00006606536,0.000063531166,0.0027582643,0.0012547685],"genre_scores_gemma":[0.23954049,0.00020749838,0.7534686,0.00029715826,0.00008127939,0.00035175282,0.0007185646,0.0004974023,0.004837324],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996172,0.00009649808,0.000025847989,0.0000680722,0.000139448,0.000052894084],"domain_scores_gemma":[0.99900025,0.0005510784,0.000075126954,0.00015920431,0.00016194617,0.000052342883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009579326,0.001472142,0.0008987105,0.00049616315,0.00038683024,0.0008620987,0.0015470418,0.0013050894,0.0042530377],"category_scores_gemma":[0.005674418,0.0007954873,0.00051061285,0.00062237703,0.00074235507,0.00192982,0.0016441997,0.002091172,0.001964164],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002591952,0.00015298219,0.0011202011,0.00017994003,0.000070339695,0.000121005374,0.00014472216,0.630366,0.0071625593,0.024607698,0.009217842,0.32659742],"study_design_scores_gemma":[0.000012035681,0.000014089144,0.000045943325,0.0000040452437,0.0000028190143,0.000013646436,0.0000069002813,0.99352276,0.0009301807,0.0050043333,0.00044091378,0.0000023105488],"about_ca_topic_score_codex":0.005367836,"about_ca_topic_score_gemma":0.011736596,"teacher_disagreement_score":0.005367836,"about_ca_system_score_codex":0.0010872481,"about_ca_system_score_gemma":0.0014998713,"threshold_uncertainty_score":0.0142278075},"labels":[],"label_agreement":null},{"id":"W2963615308","doi":"10.1609/aaai.v33i01.3301232","title":"Multi-Perspective Relevance Matching with Hierarchical ConvNets for Social Media Search","year":2019,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Ranking (information retrieval); Social media; Relevance (law); Microblogging; Information retrieval; Convolutional neural network; Phrase; Artificial intelligence; Artificial neural network; Feature (linguistics); Matching (statistics); World Wide Web","score_opus":0.10936196247753777,"score_gpt":0.3302808550666266,"score_spread":0.2209188925890888,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963615308","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12813069,0.0032327427,0.8477129,0.001121153,0.00018360185,0.00022623639,0.0012714752,0.0073721074,0.010749184],"genre_scores_gemma":[0.88020414,0.00057208055,0.10437625,0.00039823464,0.00016605007,0.0001410456,0.0013883964,0.00020342415,0.012550324],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996581,0.000074957796,0.000017850858,0.000090835005,0.000077479686,0.00008076337],"domain_scores_gemma":[0.9996265,0.0001308784,0.00006441879,0.000062031875,0.00008328201,0.000032823074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006987852,0.00089741155,0.001013501,0.0012063622,0.00040699053,0.00079297007,0.001726041,0.0012592167,0.0033873518],"category_scores_gemma":[0.0021397227,0.0004860652,0.00085779035,0.0013088599,0.00040867744,0.0020864925,0.0009142516,0.0010286847,0.0013781689],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000525386,0.00038116865,0.0029422313,0.00025907764,0.00023489424,0.0002736793,0.00014381374,0.54191506,0.008221532,0.022791568,0.015214653,0.407097],"study_design_scores_gemma":[0.000008855729,0.000028343444,0.000163018,0.0000042490537,0.000010530541,0.000017653676,0.000006632129,0.9928127,0.00053232146,0.0059863725,0.00042453682,0.000004714168],"about_ca_topic_score_codex":0.015161128,"about_ca_topic_score_gemma":0.025429888,"teacher_disagreement_score":0.015161128,"about_ca_system_score_codex":0.0014768428,"about_ca_system_score_gemma":0.0009884344,"threshold_uncertainty_score":0.030145764},"labels":[],"label_agreement":null},{"id":"W2963897632","doi":"10.18653/v1/n18-2047","title":"Strong Baselines for Simple Question Answering over Knowledge Graphs with and without Neural Networks","year":2018,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":108,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Question answering; Computer science; Simple (philosophy); Artificial neural network; Computational linguistics; Knowledge graph; Artificial intelligence; Volume (thermodynamics); Natural language processing; Data science; Epistemology; Philosophy","score_opus":0.022088721345324257,"score_gpt":0.28927770348054704,"score_spread":0.2671889821352228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963897632","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24276558,0.056698058,0.5274178,0.011129611,0.005805032,0.001516649,0.039721165,0.054773614,0.060172576],"genre_scores_gemma":[0.7555554,0.0028116505,0.16729915,0.0017758564,0.0016456835,0.0005646618,0.052109666,0.0015158533,0.016722124],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99353665,0.002248089,0.0003184694,0.002022418,0.0011750417,0.00069924956],"domain_scores_gemma":[0.98581046,0.008699181,0.00035138146,0.002994052,0.0016081466,0.0005368544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008766359,0.0028190946,0.0025153405,0.004435704,0.0021461484,0.0035868683,0.005063108,0.0057253665,0.0132502755],"category_scores_gemma":[0.030401561,0.0007850335,0.0020050893,0.0030238891,0.0012999827,0.012950433,0.0046597365,0.004279966,0.006810495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005835781,0.0019511649,0.0047096964,0.0020136947,0.001146686,0.00030386067,0.00037893184,0.1186396,0.006315636,0.029267298,0.21253908,0.6168986],"study_design_scores_gemma":[0.00034002363,0.00032515163,0.0019258949,0.00011154279,0.00028717116,0.00011398101,0.00015707478,0.8899264,0.0034463713,0.09014219,0.013168931,0.000055369448],"about_ca_topic_score_codex":0.016306547,"about_ca_topic_score_gemma":0.03100836,"teacher_disagreement_score":0.016306547,"about_ca_system_score_codex":0.0031586739,"about_ca_system_score_gemma":0.0020956385,"threshold_uncertainty_score":0.046361446},"labels":[],"label_agreement":null},{"id":"W2963903950","doi":"10.18653/v1/d16-1230","title":"How NOT To Evaluate Your Dialogue System: An Empirical Study of Unsupervised Evaluation Metrics for Dialogue Response Generation","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":916,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; McGill University","funders":"","keywords":"Computer science; Empirical research; Artificial intelligence; Machine learning; Natural language processing; Statistics; Mathematics","score_opus":0.2550579332141903,"score_gpt":0.39305640244623175,"score_spread":0.13799846923204145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963903950","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86220926,0.0022776132,0.120921336,0.0007072926,0.00027339451,0.0011573924,0.0014084174,0.0025617492,0.008483575],"genre_scores_gemma":[0.95565385,0.00014806789,0.039507657,0.00023462971,0.00007561212,0.0008589956,0.001730532,0.00068946317,0.0011011229],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8514024,0.120750315,0.006361726,0.008328507,0.01186841,0.0012886282],"domain_scores_gemma":[0.44956434,0.46450773,0.021951595,0.027448805,0.03308805,0.003439408],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0751058,0.0014015883,0.0011352077,0.0025129917,0.0016947396,0.002574272,0.0016282605,0.0015560508,0.00096870016],"category_scores_gemma":[0.34900346,0.00042931954,0.0006807425,0.0020885973,0.0019692576,0.004241251,0.0027249567,0.0028489977,0.0010157702],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005577881,0.0036316565,0.2433226,0.0038337968,0.0016563502,0.00035421152,0.021507887,0.038781475,0.02918665,0.0074942573,0.026671398,0.6179819],"study_design_scores_gemma":[0.0006611232,0.010042241,0.36798787,0.0011842129,0.00083841314,0.0016901719,0.008597869,0.4965291,0.058741756,0.019846525,0.033029255,0.0008514684],"about_ca_topic_score_codex":0.0033038289,"about_ca_topic_score_gemma":0.0036413837,"teacher_disagreement_score":0.9248942,"about_ca_system_score_codex":0.0020452822,"about_ca_system_score_gemma":0.0014927058,"threshold_uncertainty_score":0.3972022},"labels":[],"label_agreement":null},{"id":"W2963929190","doi":"10.18653/v1/k16-1028","title":"Abstractive Text Summarization using Sequence-to-sequence RNNs and Beyond","year":2016,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2209,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Automatic summarization; Sequence (biology); Computer science; Natural language processing; Artificial intelligence; Information retrieval; Biology; Genetics","score_opus":0.07923942989030704,"score_gpt":0.31578085645877263,"score_spread":0.2365414265684656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963929190","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025189819,0.0022435384,0.9627937,0.0004432987,0.00021809389,0.00013969794,0.00078640407,0.005255697,0.002929612],"genre_scores_gemma":[0.39788696,0.0021188846,0.5819323,0.0004437651,0.0005012755,0.00029524983,0.00539711,0.00074239523,0.01068205],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994697,0.00016079682,0.000041567015,0.00018790628,0.00010453793,0.000035494046],"domain_scores_gemma":[0.99858654,0.0005841265,0.00019437703,0.00022705978,0.00036170354,0.000046267764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009583443,0.0014454997,0.0008342246,0.0010754772,0.00037692397,0.0010614122,0.0012610801,0.0007744436,0.0022819084],"category_scores_gemma":[0.0035613126,0.00035643793,0.00082168414,0.0009254709,0.0003832125,0.0028380423,0.0006351213,0.0013402426,0.0016157473],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048915885,0.00022202317,0.0011414178,0.00070078194,0.00028711322,0.00024578103,0.00039776563,0.3099568,0.060585726,0.014486524,0.011625373,0.59986156],"study_design_scores_gemma":[0.000017695225,0.00013400365,0.0004162074,0.000025791913,0.000062497114,0.000050947205,0.00004031743,0.9725658,0.012437217,0.008286667,0.0059421933,0.000020619405],"about_ca_topic_score_codex":0.0043547736,"about_ca_topic_score_gemma":0.008747865,"teacher_disagreement_score":0.0043547736,"about_ca_system_score_codex":0.0006875195,"about_ca_system_score_gemma":0.00060751464,"threshold_uncertainty_score":0.008658826},"labels":[],"label_agreement":null},{"id":"W2963935808","doi":"","title":"Focused Hierarchical RNNs for Conditional Sequence Processing","year":2018,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; McGill University; Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Computer science; Security token; Recurrent neural network; Encoder; Sequence (biology); Generalization; Context (archaeology); Artificial intelligence; Embedding; Dependency (UML); Machine learning; Artificial neural network","score_opus":0.027547415482368918,"score_gpt":0.2714953414140823,"score_spread":0.2439479259317134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963935808","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011155401,0.0006944706,0.9814368,0.0001887771,0.00007385543,0.00007278694,0.0004666641,0.004050342,0.0018609427],"genre_scores_gemma":[0.42346612,0.0009216548,0.561296,0.0005538213,0.00017326996,0.0004054809,0.0033625714,0.0006052054,0.009215845],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995672,0.00012093769,0.000025053005,0.0001629634,0.0000787142,0.000045121695],"domain_scores_gemma":[0.9991304,0.0004299489,0.000079215046,0.00015047018,0.00017774853,0.00003229265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000922236,0.0012285706,0.0005848149,0.000661963,0.000266342,0.00052128785,0.0016190398,0.00087951275,0.004601006],"category_scores_gemma":[0.0028664062,0.00050100306,0.0008293509,0.0007152587,0.00041676397,0.0014825693,0.0007143343,0.0017640273,0.0018741061],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018739513,0.00014619561,0.00092684507,0.00027834444,0.00013627086,0.00013013254,0.00011542812,0.6358603,0.023407198,0.029194446,0.011779957,0.29783738],"study_design_scores_gemma":[0.000004543281,0.00001890847,0.00012258608,0.0000062732206,0.000008409192,0.000011855311,0.0000032237801,0.99035066,0.0023503783,0.0062658824,0.0008522549,0.0000049887667],"about_ca_topic_score_codex":0.011175923,"about_ca_topic_score_gemma":0.01820088,"teacher_disagreement_score":0.011175923,"about_ca_system_score_codex":0.001224873,"about_ca_system_score_gemma":0.0010794079,"threshold_uncertainty_score":0.022221744},"labels":[],"label_agreement":null},{"id":"W2964026269","doi":"10.1145/3322640.3326742","title":"Statute Law Information Retrieval and Entailment","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Paragraph; Textual entailment; Logical consequence; Statute; Information retrieval; Natural language processing; Artificial intelligence; Law; Political science; World Wide Web","score_opus":0.008123675063360826,"score_gpt":0.21202897323979214,"score_spread":0.20390529817643133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964026269","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11034833,0.0010797256,0.85451716,0.0015198404,0.00012218453,0.0011700672,0.0031457697,0.014288559,0.013808453],"genre_scores_gemma":[0.40410358,0.00040758215,0.577317,0.00058062177,0.00017354386,0.00050593563,0.011313235,0.00040233877,0.0051961127],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99419665,0.0014430908,0.0007326618,0.0010475547,0.0022333004,0.0003466564],"domain_scores_gemma":[0.992782,0.0036270148,0.00054287416,0.0010134398,0.0018650598,0.00016955793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030087326,0.0009531829,0.0013721022,0.005796846,0.0010865425,0.0024145688,0.0020435425,0.0017284177,0.0074100355],"category_scores_gemma":[0.019989608,0.000431418,0.0014674886,0.0025634677,0.00083057873,0.0047797384,0.0017391224,0.0011306928,0.003412213],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005895378,0.00086393015,0.0096539175,0.0010309814,0.00014968764,0.00056608394,0.0010176435,0.019407209,0.0676745,0.019637154,0.025940271,0.85346913],"study_design_scores_gemma":[0.00017189862,0.00045625577,0.016032895,0.00012674429,0.00016870858,0.0011849547,0.0005765155,0.7508386,0.15298392,0.041982863,0.035325874,0.00015083415],"about_ca_topic_score_codex":0.007377424,"about_ca_topic_score_gemma":0.00760012,"teacher_disagreement_score":0.0074100355,"about_ca_system_score_codex":0.0014876788,"about_ca_system_score_gemma":0.0017010957,"threshold_uncertainty_score":0.024789095},"labels":[],"label_agreement":null},{"id":"W2964035032","doi":"","title":"Towards Binary-Valued Gates for Robust LSTM Training","year":2018,"lang":"en","type":"article","venue":"International Conference on Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Generalization; Binary number; Computation; Sequence (biology); Compression (physics); Artificial intelligence; Rank (graph theory); Algorithm; Logic gate; Theoretical computer science; Arithmetic; Mathematics","score_opus":0.12127680136731257,"score_gpt":0.3393365146854237,"score_spread":0.21805971331811114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964035032","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014217505,0.0005366278,0.9799065,0.000360463,0.000073348485,0.000059179296,0.00022132728,0.0030227355,0.0016022433],"genre_scores_gemma":[0.59126365,0.0006370726,0.40022454,0.0008052499,0.00014750559,0.000380601,0.001289968,0.00058306276,0.0046682954],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99930775,0.00019310825,0.000055003944,0.00020928025,0.00014752675,0.000087394954],"domain_scores_gemma":[0.9986644,0.0008275282,0.00011049504,0.00018389699,0.00015697294,0.000056721317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010217694,0.0017176036,0.0012220173,0.0006846759,0.00031644292,0.0012694612,0.001939904,0.0018911753,0.0050180443],"category_scores_gemma":[0.005495403,0.0007378911,0.00067457417,0.0006857576,0.00093019113,0.0023413172,0.0016277465,0.0029628063,0.0018145631],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004136054,0.000110648914,0.00051262585,0.00032722123,0.000096840406,0.00018400578,0.00011242664,0.58113647,0.023091404,0.042277917,0.0052964613,0.34644043],"study_design_scores_gemma":[0.000008745359,0.000022395569,0.000025076035,0.000012040165,0.0000066502125,0.000014305592,0.000004923277,0.9832997,0.002708815,0.013465116,0.00042821706,0.0000040293594],"about_ca_topic_score_codex":0.0022647078,"about_ca_topic_score_gemma":0.0030134362,"teacher_disagreement_score":0.0050180443,"about_ca_system_score_codex":0.0010453495,"about_ca_system_score_gemma":0.0011162646,"threshold_uncertainty_score":0.016787052},"labels":[],"label_agreement":null},{"id":"W2964087600","doi":"","title":"Improved Relation Classification by Deep Recurrent Neural Networks with Data Augmentation","year":2016,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":97,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Relation (database); Artificial intelligence; Task (project management); Convolutional neural network; Abstraction; Artificial neural network; Representation (politics); Machine learning; Deep learning; Layer (electronics); Recurrent neural network; SemEval; Deep neural networks; External Data Representation; Data mining","score_opus":0.09306325500152063,"score_gpt":0.19577380823893065,"score_spread":0.10271055323741002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964087600","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27868217,0.009558408,0.66428316,0.001792799,0.00088819163,0.00034251038,0.005372268,0.026639208,0.012441395],"genre_scores_gemma":[0.7567616,0.0015488545,0.21259849,0.00053459255,0.00043229872,0.00021000694,0.015911771,0.00044706807,0.011555258],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984901,0.0003742586,0.00014155447,0.0005816761,0.0002648925,0.00014761438],"domain_scores_gemma":[0.99763274,0.0010669781,0.00024228614,0.0005957383,0.0003711384,0.000091104586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001940177,0.0021770822,0.0013256529,0.0025309692,0.00064271217,0.0014857826,0.0018068677,0.0012510068,0.0027035293],"category_scores_gemma":[0.005250442,0.0005330398,0.0018827853,0.002402172,0.00043120052,0.0047100955,0.0016285417,0.0020527255,0.0030140122],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058866973,0.0005967823,0.008192539,0.00027618746,0.00021857335,0.00026607793,0.00043439242,0.05613803,0.01489441,0.0039703394,0.023690907,0.89073306],"study_design_scores_gemma":[0.000032770156,0.00009986277,0.001990195,0.00003739419,0.00009760874,0.00009034292,0.00007175149,0.98080355,0.005164204,0.007122755,0.0044638533,0.000025688765],"about_ca_topic_score_codex":0.008531886,"about_ca_topic_score_gemma":0.012750157,"teacher_disagreement_score":0.008531886,"about_ca_system_score_codex":0.00075010303,"about_ca_system_score_gemma":0.0008639301,"threshold_uncertainty_score":0.016964495},"labels":[],"label_agreement":null},{"id":"W2964122685","doi":"10.18653/v1/p19-1153","title":"Towards Lossless Encoding of Sentences","year":2019,"lang":"","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Lossless compression; Computer science; Sentence; Encoding (memory); Embedding; Feature (linguistics); Natural language processing; Focus (optics); Artificial intelligence; Task (project management); Sequence labeling; Compression (physics); Sequence (biology); Data compression; Speech recognition; Linguistics","score_opus":0.017228668152795897,"score_gpt":0.24107585302181855,"score_spread":0.22384718486902266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964122685","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021085063,0.00071178,0.9721596,0.0006363914,0.00022632374,0.00005180041,0.000813822,0.0023208505,0.0019944562],"genre_scores_gemma":[0.35621,0.001639148,0.61912155,0.0013129088,0.0006010467,0.0004543932,0.005256681,0.0010761817,0.014328123],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9990742,0.00034597234,0.00007635703,0.00016908559,0.0002585385,0.000075808566],"domain_scores_gemma":[0.9982326,0.0006555055,0.00014769827,0.0005391542,0.00035964735,0.00006539841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008398527,0.00080260274,0.0005949671,0.00082092104,0.0002304722,0.00093435834,0.0007246296,0.00071215513,0.0030459894],"category_scores_gemma":[0.005924236,0.00026142102,0.00044534658,0.000710211,0.00058921694,0.003109876,0.0014291621,0.0015453493,0.0023569593],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075788755,0.00021560462,0.00076010387,0.00039931736,0.00006658432,0.00034986596,0.0005501464,0.07878282,0.060749386,0.07601732,0.035024285,0.7463267],"study_design_scores_gemma":[0.00006424676,0.00032667763,0.0004798994,0.000089492176,0.00004474722,0.00038090968,0.00017097278,0.8040449,0.050325196,0.11390086,0.030131688,0.000040363182],"about_ca_topic_score_codex":0.00054881966,"about_ca_topic_score_gemma":0.00089424493,"teacher_disagreement_score":0.0030459894,"about_ca_system_score_codex":0.00043199168,"about_ca_system_score_gemma":0.0005744572,"threshold_uncertainty_score":0.010189831},"labels":[],"label_agreement":null},{"id":"W2964127085","doi":"10.18653/v1/w19-3007","title":"The importance of sharing patient-generated clinical speech and language data","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Natural language processing; Speech recognition; Artificial intelligence","score_opus":0.11208722925809173,"score_gpt":0.3567917858136232,"score_spread":0.24470455655553147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964127085","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02394451,0.012655911,0.8054675,0.12171465,0.0029891501,0.0017628361,0.008200854,0.0016989843,0.021565633],"genre_scores_gemma":[0.29329547,0.011630962,0.6457276,0.015619052,0.0056082946,0.004057124,0.018363828,0.0013256153,0.004372205],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.8366076,0.1288859,0.0070800656,0.012801155,0.012825038,0.001800275],"domain_scores_gemma":[0.53961813,0.30509427,0.01555539,0.11270701,0.020692552,0.0063325902],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.13220716,0.0011531705,0.0023672937,0.006118421,0.0039400137,0.015743973,0.0058654496,0.005501502,0.006525167],"category_scores_gemma":[0.27229956,0.0015067396,0.0022840078,0.007971979,0.005771473,0.030423293,0.021997506,0.0067525203,0.0037277848],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009365625,0.00058455527,0.047460675,0.005121134,0.0017487548,0.001226112,0.022675604,0.018755762,0.006783123,0.22709829,0.07296783,0.5946416],"study_design_scores_gemma":[0.00015714286,0.00035035817,0.013799826,0.0025887375,0.00047982187,0.0016921298,0.012397997,0.017346695,0.004426538,0.6507279,0.29566863,0.0003641063],"about_ca_topic_score_codex":0.0038101897,"about_ca_topic_score_gemma":0.0033484006,"teacher_disagreement_score":0.99413455,"about_ca_system_score_codex":0.003186581,"about_ca_system_score_gemma":0.012513134,"threshold_uncertainty_score":0.6991866},"labels":[],"label_agreement":null},{"id":"W2964166731","doi":"10.18653/v1/p19-1128","title":"Graph Neural Networks with Generated Parameters for Relation Extraction","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":176,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Relationship extraction; Artificial neural network; Artificial intelligence; Graph; Relation (database); Generator (circuit theory); Machine learning; Natural language processing; Data mining; Theoretical computer science; Power (physics)","score_opus":0.04244054077691106,"score_gpt":0.26377833464473355,"score_spread":0.22133779386782249,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964166731","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04254036,0.0021195926,0.88757676,0.001043413,0.00042083862,0.0004041733,0.015421461,0.037770506,0.0127029065],"genre_scores_gemma":[0.48679245,0.001214384,0.45940915,0.00036161934,0.00023997713,0.0006422808,0.032350652,0.0029236316,0.016065815],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99922574,0.00020412797,0.000052145213,0.0003360219,0.00012272024,0.000059251688],"domain_scores_gemma":[0.9982761,0.00095278147,0.00006834014,0.00043893614,0.00023124133,0.000032640888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009021169,0.0013521565,0.00063030625,0.002589645,0.0005226142,0.0014294261,0.001458127,0.0020145003,0.018042274],"category_scores_gemma":[0.008475043,0.000671794,0.0011032069,0.003118905,0.0003698392,0.002864921,0.0010426367,0.0023079694,0.0103782555],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060953666,0.00022582752,0.00245588,0.00047228386,0.00022117552,0.00017475328,0.00016475683,0.14011027,0.011147803,0.013642282,0.049918484,0.78085685],"study_design_scores_gemma":[0.000046095258,0.000037016263,0.0008387576,0.000056071684,0.00008028421,0.00007322751,0.00004183921,0.95743895,0.006182518,0.025432842,0.009753826,0.000018607994],"about_ca_topic_score_codex":0.009643019,"about_ca_topic_score_gemma":0.014764577,"teacher_disagreement_score":0.018042274,"about_ca_system_score_codex":0.0010713413,"about_ca_system_score_gemma":0.0010328896,"threshold_uncertainty_score":0.06035739},"labels":[],"label_agreement":null},{"id":"W2965121914","doi":"","title":"Hybrid Deep Neural Networks to Predict Socio-Moral Reasoning Skills.","year":2019,"lang":"en","type":"article","venue":"Espace ÉTS (ETS)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec; Université du Québec à Montréal","funders":"","keywords":"Artificial neural network; Computer science; Artificial intelligence; Deep neural networks; Machine learning","score_opus":0.008193866942364045,"score_gpt":0.2282034845638554,"score_spread":0.22000961762149135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965121914","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8016484,0.0032849258,0.17519626,0.001901006,0.00042355774,0.00010883591,0.0031078816,0.001034372,0.013294857],"genre_scores_gemma":[0.9788538,0.00027070922,0.013503749,0.00008407865,0.0000620641,0.000047115205,0.0019790318,0.000029383003,0.005169999],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997845,0.00009226603,0.000008008727,0.000048312028,0.000029843248,0.00003712374],"domain_scores_gemma":[0.99903727,0.000607668,0.00006689076,0.0000536317,0.00015960744,0.00007498572],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010856801,0.0004991464,0.00027765348,0.00074704847,0.00019452523,0.00073867943,0.000600971,0.00065670104,0.0033307152],"category_scores_gemma":[0.0028695657,0.00018388186,0.00051373633,0.00049271935,0.00014310176,0.00089685724,0.00056384725,0.0016791406,0.0011016494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001228006,0.0018705009,0.07070564,0.0002287738,0.00064908474,0.00014057684,0.00047018263,0.3492593,0.006283485,0.010725849,0.0337016,0.5247371],"study_design_scores_gemma":[0.000013149415,0.00003658549,0.0056090555,0.000024156569,0.000025617994,0.000008730227,0.00004089132,0.9872744,0.00062300876,0.0055375597,0.0008009099,0.0000059760855],"about_ca_topic_score_codex":0.0071397494,"about_ca_topic_score_gemma":0.013173426,"teacher_disagreement_score":0.0071397494,"about_ca_system_score_codex":0.00071801234,"about_ca_system_score_gemma":0.0004985796,"threshold_uncertainty_score":0.014196396},"labels":[],"label_agreement":null},{"id":"W2965233127","doi":"10.24963/ijcai.2019/731","title":"Modeling Noisy Hierarchical Types in Fine-Grained Entity Typing: A Content-Based Weighting Approach","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; University of Ottawa","funders":"National Key Research and Development Program of China; Beijing Advanced Innovation Center for Big Data and Brain Computing; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China","keywords":"Computer science; Weighting; Artificial intelligence; Benchmark (surveying); Embedding; Schema (genetic algorithms); Process (computing); Noisy data; Sentence; Set (abstract data type); Data mining; Machine learning; Natural language processing","score_opus":0.045829028796194674,"score_gpt":0.24012107199333468,"score_spread":0.19429204319714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965233127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02116673,0.00022418523,0.97725266,0.00010409805,0.000038531747,0.00006784566,0.00019039458,0.0005339519,0.0004215951],"genre_scores_gemma":[0.4504804,0.00049006165,0.5398435,0.00027134523,0.00017984865,0.00034590857,0.002065034,0.00058899843,0.0057348926],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99816126,0.0006013972,0.00015046363,0.00058140076,0.00035751727,0.00014781128],"domain_scores_gemma":[0.99388933,0.0031637796,0.0006486718,0.0011676656,0.0009322169,0.00019830417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034121824,0.0011358659,0.0013554334,0.0030430558,0.00079877826,0.0014117756,0.002784884,0.0019417215,0.0013950499],"category_scores_gemma":[0.012300097,0.0007308327,0.0012321145,0.0038036576,0.00084213226,0.0055712587,0.0021696608,0.0023060355,0.0010960668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005362451,0.000582065,0.0193784,0.000514654,0.00028392376,0.00053672196,0.001833786,0.26890844,0.028091263,0.050769307,0.011174327,0.6173908],"study_design_scores_gemma":[0.000013514251,0.000047981815,0.0012739691,0.000029087267,0.000049562226,0.00013567641,0.0001073295,0.9617507,0.0039858283,0.029897874,0.0026811785,0.00002736731],"about_ca_topic_score_codex":0.004013664,"about_ca_topic_score_gemma":0.009351993,"teacher_disagreement_score":0.004013664,"about_ca_system_score_codex":0.0008825252,"about_ca_system_score_gemma":0.00090368045,"threshold_uncertainty_score":0.018045604},"labels":[],"label_agreement":null},{"id":"W2965484644","doi":"10.48550/arxiv.1907.11843","title":"Analyzing Linguistic Complexity and Scientific Impact","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Citation; Proxy (statistics); Scientific literature; Scientific writing; Context (archaeology); Scientific communication; Linguistics; Linguistic context; Linguistic sequence complexity; Linguistic analysis; Psychology; Computer science; Library science; History; Biology; Philosophy","score_opus":0.1372479509848239,"score_gpt":0.2273120980192345,"score_spread":0.0900641470344106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965484644","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.991846,0.0006993027,0.0025473048,0.00029499654,0.000018769486,0.00004163752,0.0008456401,0.000037526497,0.003668829],"genre_scores_gemma":[0.9971282,0.00018067092,0.0013395877,0.00001610946,0.00007289949,0.000044954322,0.00093253964,0.000013495741,0.0002715572],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9954026,0.001716199,0.00054869964,0.0004218706,0.0016197202,0.00029088854],"domain_scores_gemma":[0.854618,0.11706144,0.01800898,0.0030650739,0.0052563883,0.0019901525],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0059038107,0.0005136986,0.0007452385,0.018319463,0.0009724893,0.0045713596,0.00042776662,0.0006411432,0.0027565486],"category_scores_gemma":[0.07213692,0.00021703992,0.00084567,0.0162916,0.0010538528,0.0024448517,0.0026750013,0.00092322286,0.0004442726],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041060155,0.00022085494,0.93426037,0.0004046832,0.00072142517,0.00026036825,0.0039261472,0.006154673,0.0020489944,0.003125471,0.0009538164,0.047512688],"study_design_scores_gemma":[0.000025156696,0.00011216269,0.9656019,0.00008376878,0.0002394217,0.00019346534,0.0024877654,0.019802906,0.0009029754,0.008056365,0.0024315934,0.00006258185],"about_ca_topic_score_codex":0.002548012,"about_ca_topic_score_gemma":0.0017155135,"teacher_disagreement_score":0.99409616,"about_ca_system_score_codex":0.0010155858,"about_ca_system_score_gemma":0.0008941453,"threshold_uncertainty_score":0.031222701},"labels":[],"label_agreement":null},{"id":"W2965587560","doi":"","title":"End-to-end Neural Information Retrieval","year":2019,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo","keywords":"Information retrieval; End-to-end principle; Computer science; Artificial neural network; Artificial intelligence","score_opus":0.010519293261457988,"score_gpt":0.19852714659480175,"score_spread":0.18800785333334377,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965587560","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06319491,0.0053708255,0.83921486,0.001898539,0.0006859816,0.0012247904,0.00798677,0.057093795,0.023329541],"genre_scores_gemma":[0.35021704,0.001931757,0.5550207,0.00140163,0.0003913112,0.00067665754,0.021322737,0.0010588088,0.06797927],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986436,0.0002106406,0.0001180054,0.0004929879,0.00036711508,0.00016767153],"domain_scores_gemma":[0.9985875,0.00033937336,0.00008509271,0.0004206993,0.0005099227,0.00005734776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014983303,0.001933616,0.0014884553,0.0018809858,0.0009816955,0.0018897195,0.003514538,0.0021760168,0.016115844],"category_scores_gemma":[0.005635909,0.000556455,0.0009869561,0.0018229011,0.0005513203,0.0039175036,0.0017707868,0.0018294454,0.014299902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061118975,0.0006833573,0.001229247,0.00051702786,0.00019953359,0.00020673,0.00010482972,0.067247376,0.016992254,0.00433728,0.069461934,0.83840925],"study_design_scores_gemma":[0.00007274264,0.00023772364,0.00081534806,0.000031350253,0.00006757146,0.00014457869,0.000078744655,0.95117456,0.023872077,0.010428267,0.013027971,0.000049020335],"about_ca_topic_score_codex":0.017853254,"about_ca_topic_score_gemma":0.03581155,"teacher_disagreement_score":0.017853254,"about_ca_system_score_codex":0.0020322825,"about_ca_system_score_gemma":0.0018275804,"threshold_uncertainty_score":0.053912878},"labels":[],"label_agreement":null},{"id":"W2965774554","doi":"10.1109/access.2019.2933354","title":"Syntactic, Semantic and Sentiment Analysis: The Joint Effect on Automated Essay Evaluation","year":2019,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Natural Sciences and Engineering Research Council of Canada; Lakehead University","keywords":"Computer science; Artificial intelligence; Natural language processing; Semantic similarity; Sentiment analysis; Syntax; Graph; Similarity (geometry); Information retrieval; Machine learning; Theoretical computer science","score_opus":0.0303102827429724,"score_gpt":0.3266139024221374,"score_spread":0.29630361967916496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965774554","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3196807,0.0016348133,0.62572896,0.0016938958,0.0004558315,0.00066862075,0.0011318136,0.017060418,0.031945035],"genre_scores_gemma":[0.808731,0.00030005042,0.18197995,0.00017957337,0.00026174157,0.00018556771,0.0009216995,0.0005138006,0.0069265184],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99071455,0.005119281,0.00062732695,0.0010594724,0.002195238,0.0002840284],"domain_scores_gemma":[0.97551936,0.012673227,0.0019193584,0.0019495763,0.007412149,0.000526352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008405777,0.001150261,0.0012449133,0.004616833,0.0008450562,0.0030472698,0.0006929269,0.0008208577,0.0051519093],"category_scores_gemma":[0.029558035,0.00043285423,0.00076340104,0.0021242292,0.0005504726,0.003137971,0.002068356,0.00089754205,0.003586379],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047402884,0.0002550513,0.017437285,0.00016169924,0.00015140371,0.00009391067,0.00034215656,0.005196404,0.026004413,0.001286299,0.008284171,0.94031316],"study_design_scores_gemma":[0.00009722046,0.0006206283,0.0787348,0.00011740475,0.0002532209,0.00039535516,0.00066927384,0.8456721,0.04354907,0.013205826,0.016514756,0.0001702977],"about_ca_topic_score_codex":0.0028342973,"about_ca_topic_score_gemma":0.0039999476,"teacher_disagreement_score":0.008405777,"about_ca_system_score_codex":0.000570145,"about_ca_system_score_gemma":0.0009974051,"threshold_uncertainty_score":0.044454515},"labels":[],"label_agreement":null},{"id":"W2965914381","doi":"10.1109/visual.2019.8933744","title":"Visualizing RNN States with Predictive Semantic Encodings","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Intuition; Recurrent neural network; Artificial intelligence; Visualization; Natural language; Semantics (computer science); Natural language processing; Task (project management); Encoding (memory); Artificial neural network; Programming language; Cognitive science","score_opus":0.020631734389863935,"score_gpt":0.26448334919274974,"score_spread":0.2438516148028858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2965914381","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056925,0.0004396912,0.9171455,0.0008401696,0.00021504781,0.000063637395,0.0025188304,0.012509926,0.009342333],"genre_scores_gemma":[0.66107434,0.0006061561,0.3298615,0.00013082867,0.00006256484,0.0001627961,0.001969537,0.0011707036,0.004961531],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998859,0.000038091777,0.0000073817528,0.000025335956,0.000030537685,0.000012730852],"domain_scores_gemma":[0.9995933,0.00019816498,0.000044114357,0.000058488924,0.000082820115,0.000023162562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00033297358,0.0008307556,0.00024533345,0.0006951841,0.00018390676,0.0010654806,0.00048135757,0.00053433905,0.008170189],"category_scores_gemma":[0.0018736165,0.00019998639,0.00044124806,0.00049400603,0.00031285433,0.0013062658,0.0008530156,0.00078400224,0.0010115905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009991105,0.00015023914,0.0029779456,0.0008509265,0.00010050168,0.0009402129,0.0027590876,0.44318464,0.091154516,0.17105883,0.032425616,0.25339836],"study_design_scores_gemma":[0.000032287964,0.00005342284,0.00083983794,0.000078109195,0.000019001165,0.00011010616,0.000186519,0.9159236,0.01492885,0.05360774,0.014187745,0.000032832348],"about_ca_topic_score_codex":0.0025089323,"about_ca_topic_score_gemma":0.0025831158,"teacher_disagreement_score":0.008170189,"about_ca_system_score_codex":0.00040177983,"about_ca_system_score_gemma":0.00034841473,"threshold_uncertainty_score":0.027332008},"labels":[],"label_agreement":null},{"id":"W2966132766","doi":"10.1609/aaai.v33i01.33019987","title":"A Multi-Task Learning Framework for Abstractive Text Summarization","year":2019,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Syntax; Natural language processing; Categorization; Task (project management); Artificial intelligence; Encoder","score_opus":0.08044737931383913,"score_gpt":0.31492123089233054,"score_spread":0.2344738515784914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966132766","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045569357,0.0005533742,0.99011505,0.00034671105,0.00011398478,0.00012291182,0.0002913775,0.0030888051,0.00081084133],"genre_scores_gemma":[0.18572082,0.0006064523,0.800123,0.0006787547,0.000603762,0.0007970525,0.004008216,0.00048197422,0.00697991],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837613,0.00063462846,0.00010583619,0.00045464607,0.00030058512,0.0001281505],"domain_scores_gemma":[0.9980725,0.00075243245,0.00019758478,0.0003151166,0.0005222099,0.00014009082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024715355,0.0020865847,0.0014884936,0.001585217,0.000715645,0.0013816949,0.0028002947,0.0021308549,0.0038639249],"category_scores_gemma":[0.0051809466,0.00049561437,0.0014885085,0.001553275,0.00067903963,0.0036179055,0.0020992944,0.003065494,0.0028439534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004858913,0.00061373465,0.0007119018,0.000568977,0.0002470061,0.00023848975,0.00036797166,0.1739741,0.029579943,0.012146664,0.022187704,0.75887775],"study_design_scores_gemma":[0.000039859908,0.00016375339,0.00022069729,0.000015409583,0.00003812138,0.000042333,0.000044754277,0.97644275,0.0058279643,0.012823509,0.0043173484,0.000023460023],"about_ca_topic_score_codex":0.0026967342,"about_ca_topic_score_gemma":0.0039830985,"teacher_disagreement_score":0.0038639249,"about_ca_system_score_codex":0.0008779308,"about_ca_system_score_gemma":0.0013623124,"threshold_uncertainty_score":0.013070881},"labels":[],"label_agreement":null},{"id":"W2966194045","doi":"10.1007/s42113-019-00046-x","title":"Correction to: An Instance Theory of Semantic Memory","year":2019,"lang":"en","type":"article","venue":"Computational Brain & Behavior","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Semantic memory; Natural language processing; Cognitive science; Artificial intelligence; Semantic theory of truth; Cognitive psychology; Psychology; Linguistics; Philosophy; Neuroscience; Cognition","score_opus":0.02157729932495907,"score_gpt":0.27988757523807395,"score_spread":0.2583102759131149,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966194045","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00054362265,0.00058655697,0.002051907,0.09737276,0.89361227,0.0000294464,0.0022134702,0.00077972526,0.0028102843],"genre_scores_gemma":[0.08727977,0.004352533,0.011549956,0.094266854,0.6087717,0.00037688692,0.004752734,0.0019631474,0.18668643],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972299,0.0004163569,0.00062878506,0.00060152746,0.0007795015,0.0003439998],"domain_scores_gemma":[0.96312654,0.0106297415,0.0016909965,0.0042611375,0.018929001,0.0013626455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025950945,0.0019071003,0.002625497,0.0038668755,0.0032033601,0.004252876,0.0044843405,0.008492427,0.10586399],"category_scores_gemma":[0.08037617,0.0010153735,0.0016375467,0.003276301,0.002843491,0.0031700446,0.0020705971,0.009612375,0.037534665],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000095960386,0.000010440991,0.00017227481,0.00014728658,0.000034116478,0.0005352388,0.00005825106,0.0000923741,0.00007812245,0.0035985452,0.98624945,0.008927976],"study_design_scores_gemma":[0.0001900165,0.00004495383,0.0023149746,0.00032350284,0.00010849585,0.0027021193,0.00027192742,0.0019290773,0.0009495209,0.02200075,0.9690465,0.000118113065],"about_ca_topic_score_codex":0.008215667,"about_ca_topic_score_gemma":0.008605335,"teacher_disagreement_score":0.10586399,"about_ca_system_score_codex":0.0033073623,"about_ca_system_score_gemma":0.0030937525,"threshold_uncertainty_score":0.3541503},"labels":[],"label_agreement":null},{"id":"W2966345719","doi":"10.48550/arxiv.1804.08053","title":"Learning Sentence Embeddings for Coherence Modelling and Beyond","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Coherence (philosophical gambling strategy); Heuristics; Sentence; Artificial intelligence; Structuring; Embedding; Natural language processing; Task (project management)","score_opus":0.0829964850624605,"score_gpt":0.2002490946187492,"score_spread":0.1172526095562887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966345719","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02752734,0.00053074077,0.9675156,0.00035689195,0.00008105431,0.000050225637,0.0006089106,0.0021818127,0.0011474588],"genre_scores_gemma":[0.5734907,0.00059286295,0.41735178,0.000295165,0.00023783697,0.00023535275,0.003745605,0.0006524703,0.0033981435],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99932337,0.00028980506,0.00005935538,0.00019601115,0.00008761889,0.000043985467],"domain_scores_gemma":[0.9971867,0.0015631113,0.00036366557,0.00043121018,0.00037136945,0.00008398386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010605257,0.0010055136,0.0004622571,0.0011243653,0.00039058394,0.0009939077,0.0008629991,0.0008941001,0.002530921],"category_scores_gemma":[0.008517977,0.00034697927,0.0006106942,0.0010476869,0.0005185266,0.0037597993,0.0013249763,0.0016752436,0.0011945599],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005432129,0.00033350356,0.006830417,0.00076511997,0.0001771138,0.00041468276,0.0019045688,0.15302937,0.039405767,0.06483123,0.022239719,0.70952535],"study_design_scores_gemma":[0.000023655051,0.00012359327,0.0008828051,0.000051414536,0.00003190855,0.00009935707,0.00015192639,0.9278081,0.006974196,0.056205492,0.0076220958,0.000025570573],"about_ca_topic_score_codex":0.0012099976,"about_ca_topic_score_gemma":0.00231446,"teacher_disagreement_score":0.002530921,"about_ca_system_score_codex":0.0003421313,"about_ca_system_score_gemma":0.00054261094,"threshold_uncertainty_score":0.00846684},"labels":[],"label_agreement":null},{"id":"W2966381300","doi":"10.48550/arxiv.1907.12697","title":"Dual-FOFE-net Neural Models for Entity Linking with PageRank","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Artificial intelligence; Recurrent neural network; Convolutional neural network; ENCODE; Task (project management); Forgetting; Dual (grammatical number); PageRank; Artificial neural network; Ranking (information retrieval); Cluster analysis; Machine learning; Theoretical computer science","score_opus":0.09641080250186186,"score_gpt":0.18676043684463028,"score_spread":0.09034963434276842,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966381300","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030881904,0.0012534091,0.95352536,0.0004899121,0.00020754196,0.00009723067,0.0010427077,0.006502626,0.0059992583],"genre_scores_gemma":[0.61885864,0.0010657472,0.34613705,0.0003670777,0.00029093842,0.00030572223,0.005346731,0.00044394511,0.027184058],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99962986,0.00008985559,0.000025680692,0.00014022086,0.0000711973,0.000043119922],"domain_scores_gemma":[0.99916816,0.00037533639,0.0000916659,0.00016215065,0.00017068833,0.000031978052],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009737752,0.0011115691,0.00069620047,0.001634861,0.00046732154,0.0011910225,0.0020681818,0.0013335661,0.0052356715],"category_scores_gemma":[0.003240613,0.00048216534,0.0007397044,0.0016390675,0.00044494413,0.0032672023,0.0008841146,0.0016333036,0.0024956937],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016454181,0.00016843129,0.0011443003,0.0001341172,0.0001035116,0.00015929728,0.000077426856,0.6029678,0.002805648,0.020918397,0.010631737,0.36072478],"study_design_scores_gemma":[0.000004769704,0.000013057995,0.00007053123,0.000005014071,0.0000079749325,0.000015795202,0.0000035090357,0.9903301,0.00081868673,0.007813188,0.0009124456,0.0000048660795],"about_ca_topic_score_codex":0.00638452,"about_ca_topic_score_gemma":0.0130891595,"teacher_disagreement_score":0.00638452,"about_ca_system_score_codex":0.0010171041,"about_ca_system_score_gemma":0.0007291617,"threshold_uncertainty_score":0.017515063},"labels":[],"label_agreement":null},{"id":"W2966574105","doi":"","title":"Understanding Posterior Collapse in Generative Latent Variable Models","year":2019,"lang":"en","type":"article","venue":"International Conference on Learning Representations","topic":"Topic Modeling","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Latent variable; Generative grammar; Computer science; Latent variable model; Variable (mathematics); Artificial intelligence; Mathematics","score_opus":0.17610585578987872,"score_gpt":0.3340231839156236,"score_spread":0.15791732812574485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966574105","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026293816,0.00042612743,0.9696721,0.0009639701,0.00005246225,0.00002852198,0.00029652976,0.0006284046,0.0016380522],"genre_scores_gemma":[0.79550093,0.0011477546,0.19313419,0.00048489266,0.0003263286,0.00021785426,0.0026999863,0.0009345081,0.0055535967],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99806124,0.0010970305,0.000088504974,0.00032691815,0.00026905147,0.00015728471],"domain_scores_gemma":[0.9714516,0.025526332,0.0007290779,0.0011792216,0.00074367394,0.00036996332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005669981,0.00092152884,0.0014810051,0.0017521456,0.00092208746,0.003908772,0.0021567107,0.0024667273,0.0052564545],"category_scores_gemma":[0.041136336,0.0015761872,0.0017677887,0.0015280872,0.0018711595,0.0054640328,0.003965156,0.0048292265,0.0008923541],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020082282,0.000109780834,0.003949687,0.0002600152,0.00016646988,0.00032094363,0.001566466,0.2964935,0.0016998299,0.62887394,0.006453747,0.05990483],"study_design_scores_gemma":[0.000013573969,0.000008339178,0.0003346339,0.000027562432,0.000017138696,0.00003420374,0.000053069885,0.7263889,0.00022925307,0.27205533,0.00082567986,0.000012356141],"about_ca_topic_score_codex":0.007598369,"about_ca_topic_score_gemma":0.0074245064,"teacher_disagreement_score":0.007598369,"about_ca_system_score_codex":0.0016139031,"about_ca_system_score_gemma":0.0011505581,"threshold_uncertainty_score":0.029986084},"labels":[],"label_agreement":null},{"id":"W2966610483","doi":"10.48550/arxiv.1907.12009","title":"Representation Degeneration Problem in Training Natural Language\\n Generation Models","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":105,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Regularization (linguistics); Artificial intelligence; Machine translation; Representation (politics); Maximization; Natural language processing; Language model; Natural language; Natural language understanding; Tying; Machine learning; Mathematical optimization; Mathematics","score_opus":0.16495969665061164,"score_gpt":0.21955934847399716,"score_spread":0.054599651823385525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966610483","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07138364,0.0006291447,0.9236191,0.0010735845,0.000060087277,0.000076055796,0.0001400697,0.0013673303,0.0016509503],"genre_scores_gemma":[0.7471437,0.00044595686,0.24448375,0.0009408613,0.00016442433,0.00036807306,0.0013830048,0.0004391177,0.004631225],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976271,0.0012390928,0.00014955444,0.0005726489,0.00025319165,0.00015835871],"domain_scores_gemma":[0.99130964,0.006590974,0.0004382056,0.00091692066,0.00057326455,0.00017099713],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060526836,0.0012141799,0.0014538241,0.00084060774,0.00077641354,0.0013471036,0.0022096457,0.002164594,0.0019497325],"category_scores_gemma":[0.021818444,0.0010138298,0.001010145,0.0011503033,0.0015276094,0.0041996967,0.0026653758,0.003566923,0.0010084411],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027545265,0.00023823463,0.0037846996,0.0001982787,0.00011539935,0.00021109194,0.00038439495,0.76797664,0.0042602485,0.023288751,0.0055295518,0.19373734],"study_design_scores_gemma":[0.000012445986,0.000023274859,0.00008107114,0.0000074812383,0.0000049495943,0.000020022206,0.000014248628,0.9887587,0.0008748999,0.009878499,0.00031981498,0.00000463823],"about_ca_topic_score_codex":0.0039051257,"about_ca_topic_score_gemma":0.005250101,"teacher_disagreement_score":0.0060526836,"about_ca_system_score_codex":0.0014439003,"about_ca_system_score_gemma":0.0013941852,"threshold_uncertainty_score":0.03201008},"labels":[],"label_agreement":null},{"id":"W2967813257","doi":"10.1007/s10791-019-09361-0","title":"Evaluating sentence-level relevance feedback for high-recall information retrieval","year":2019,"lang":"en","type":"article","venue":"Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Relevance feedback; Relevance (law); Recall; Computer science; Sentence; Baseline (sea); Information retrieval; Precision and recall; Natural language processing; Artificial intelligence; Cognitive psychology; Psychology; Image retrieval","score_opus":0.046889037076321534,"score_gpt":0.292025163440755,"score_spread":0.2451361263644335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967813257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75100553,0.015410303,0.21872827,0.0008884288,0.0008114804,0.0008568126,0.0013804425,0.0057659834,0.005152695],"genre_scores_gemma":[0.9170604,0.0011110565,0.074843355,0.00022563324,0.00047239405,0.00019080921,0.0027117739,0.00019716885,0.0031874273],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99647385,0.0017783897,0.00027790756,0.00037794345,0.0009265538,0.00016531562],"domain_scores_gemma":[0.9842606,0.012574472,0.00046355146,0.0003900566,0.0020153825,0.00029586742],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059652943,0.0011963488,0.0015838961,0.001913242,0.00059730833,0.0012882733,0.0008333756,0.0018683671,0.0033165875],"category_scores_gemma":[0.02758873,0.0003225321,0.00069064717,0.00072589726,0.0003188745,0.0015665174,0.0007450907,0.0008624804,0.0015494417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010543714,0.002388965,0.009490937,0.0026289432,0.00092818943,0.00043279587,0.0005421565,0.044052932,0.14398175,0.0010948478,0.01943057,0.7644842],"study_design_scores_gemma":[0.00087746,0.004431417,0.018350787,0.00011072237,0.0012011834,0.00056732364,0.00024981386,0.9054238,0.063661344,0.00205599,0.00294272,0.00012732879],"about_ca_topic_score_codex":0.0037368739,"about_ca_topic_score_gemma":0.0049501276,"teacher_disagreement_score":0.0059652943,"about_ca_system_score_codex":0.00070017274,"about_ca_system_score_gemma":0.0011364227,"threshold_uncertainty_score":0.031547844},"labels":[],"label_agreement":null},{"id":"W2968398601","doi":"10.18653/v1/d19-1458","title":"CLUTRR: A Diagnostic Benchmark for Inductive Reasoning from Text","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Mila - Quebec Artificial Intelligence Institute; McGill University","funders":"Samsung; Samsung Advanced Institute of Technology; Canadian Institute for Advanced Research","keywords":"Robustness (evolution); Computer science; Natural language understanding; Inductive logic programming; Artificial intelligence; Suite; Benchmark (surveying); Generalization; Machine learning; Theoretical computer science; Natural language; Mathematics","score_opus":0.027357898001301048,"score_gpt":0.2664366112061615,"score_spread":0.23907871320486046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2968398601","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1080913,0.0104430085,0.46801734,0.0068018553,0.0017793807,0.0025612602,0.10353783,0.25046337,0.048304625],"genre_scores_gemma":[0.24193579,0.0012969582,0.6049912,0.000997317,0.0002612272,0.0013755648,0.13547558,0.0066862553,0.006980083],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9885283,0.0044398536,0.0014357378,0.0018363998,0.0032211477,0.0005384171],"domain_scores_gemma":[0.94256365,0.043988995,0.001350498,0.00588793,0.0051492266,0.0010597274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007734243,0.0025310018,0.0010735339,0.0073505635,0.0015055463,0.0034297102,0.0073481216,0.003881194,0.013945911],"category_scores_gemma":[0.06554406,0.00092104997,0.0018469249,0.004802222,0.001697244,0.0060729957,0.004626847,0.0029663565,0.008088778],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024547821,0.0018295068,0.0078017535,0.0067592347,0.0005282257,0.0012522254,0.0011041858,0.060994618,0.007050586,0.039038245,0.35744137,0.51374525],"study_design_scores_gemma":[0.0016435771,0.0005708609,0.0030752146,0.00075551355,0.00024230455,0.0011346946,0.001183265,0.69132453,0.033560645,0.11914232,0.14721537,0.00015173577],"about_ca_topic_score_codex":0.008758047,"about_ca_topic_score_gemma":0.010765239,"teacher_disagreement_score":0.013945911,"about_ca_system_score_codex":0.0023687887,"about_ca_system_score_gemma":0.0035884685,"threshold_uncertainty_score":0.046653688},"labels":[],"label_agreement":null},{"id":"W2970008578","doi":"10.18653/v1/d19-1235","title":"Predicting Discourse Structure using Distant Supervision from Sentiment","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Huawei Technologies","keywords":"Computer science; Natural language processing; Joint (building); Artificial intelligence; Sentiment analysis; Natural (archaeology); Linguistics; History; Engineering; Philosophy","score_opus":0.015138556839577332,"score_gpt":0.2579810429993879,"score_spread":0.24284248615981058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970008578","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7727838,0.011505267,0.16961803,0.004537622,0.0014742127,0.00032476644,0.010820925,0.0057665408,0.02316878],"genre_scores_gemma":[0.9572897,0.0009805416,0.02700198,0.00013877313,0.0006693801,0.00008794188,0.009311702,0.00018735311,0.0043325624],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925965,0.00030893693,0.000034119046,0.00020927377,0.000121137426,0.00006695633],"domain_scores_gemma":[0.9958484,0.0024827567,0.00034298637,0.00025967887,0.0008526951,0.00021348443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018170418,0.0009728721,0.0006416506,0.002119001,0.00061471737,0.0012990605,0.0006407556,0.0008570019,0.0022334866],"category_scores_gemma":[0.008536699,0.00038005458,0.00050433376,0.0011898687,0.00031176995,0.001972896,0.001015838,0.0015296158,0.0030605623],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027496833,0.0010898599,0.06687627,0.00077884994,0.0005724157,0.00045804394,0.0013740184,0.028685138,0.048847347,0.0038503138,0.09484731,0.7498708],"study_design_scores_gemma":[0.00013358274,0.00041250145,0.034593366,0.00011841572,0.00023898103,0.0001057295,0.00055777445,0.9230463,0.016750468,0.010808033,0.013178727,0.000056067165],"about_ca_topic_score_codex":0.00408973,"about_ca_topic_score_gemma":0.008447643,"teacher_disagreement_score":0.00408973,"about_ca_system_score_codex":0.0005357182,"about_ca_system_score_gemma":0.00037516473,"threshold_uncertainty_score":0.00960958},"labels":[],"label_agreement":null},{"id":"W2970060558","doi":"","title":"Ordered Memory","year":2019,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Treebank; Inference; Recurrent neural network; Artificial intelligence; Task (project management); Representation (politics); Tree (set theory); Machine learning; Theoretical computer science; Artificial neural network; Natural language processing; Annotation","score_opus":0.01587798405192172,"score_gpt":0.23179901244507461,"score_spread":0.2159210283931529,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970060558","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0775162,0.0018305009,0.88585526,0.0008905049,0.0005394128,0.00010203678,0.0017746272,0.008005511,0.023485912],"genre_scores_gemma":[0.8224103,0.0011481574,0.13297029,0.00060763187,0.00017793765,0.00017323172,0.0026735957,0.00047610953,0.039362784],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998055,0.000020336582,0.00001457197,0.00008152783,0.000037760317,0.000040359744],"domain_scores_gemma":[0.99959296,0.0000988243,0.000045745055,0.00013220761,0.00010084496,0.000029421035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00029049077,0.0008725471,0.0005939071,0.00043263892,0.00037591098,0.001120013,0.002024236,0.00076142274,0.012344557],"category_scores_gemma":[0.0014011654,0.00034434695,0.0006589367,0.0005831837,0.0005418566,0.0032532753,0.0009600858,0.0010570241,0.0034052623],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059501687,0.00025512427,0.0035428084,0.00045034537,0.00018559891,0.0004910905,0.00027796472,0.18200111,0.027303826,0.1264789,0.02889665,0.6295216],"study_design_scores_gemma":[0.000028865796,0.00013387846,0.00053250813,0.000034637895,0.000066230416,0.00016384965,0.00004529337,0.8892463,0.0119904755,0.08575843,0.011967211,0.000032358847],"about_ca_topic_score_codex":0.004372366,"about_ca_topic_score_gemma":0.009343578,"teacher_disagreement_score":0.012344557,"about_ca_system_score_codex":0.0005811661,"about_ca_system_score_gemma":0.00080269907,"threshold_uncertainty_score":0.04129672},"labels":[],"label_agreement":null},{"id":"W2970161670","doi":"10.18653/v1/w19-5025","title":"Can Character Embeddings Improve Cause-of-Death Classification for Verbal Autopsy Narratives?","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Verbal autopsy; Character (mathematics); Narrative; Coding (social sciences); Computer science; Natural language processing; Autopsy; Artificial intelligence; Word (group theory); Cause of death; Speech recognition; Linguistics; Medicine; Pathology; Mathematics; Statistics","score_opus":0.03445423808597654,"score_gpt":0.28287964064471305,"score_spread":0.24842540255873652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970161670","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5642365,0.00520352,0.39929593,0.0084952265,0.0016166708,0.00034657874,0.0054938165,0.005344147,0.009967589],"genre_scores_gemma":[0.9210026,0.0007262569,0.06916663,0.0004468075,0.00032961863,0.00012366984,0.0048279334,0.00017662237,0.0031997755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99892044,0.00048612943,0.000078360696,0.00029415727,0.00010161616,0.00011938419],"domain_scores_gemma":[0.993654,0.0039837933,0.0005707165,0.0007389276,0.0008400662,0.00021250432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030179266,0.0016712578,0.0007110152,0.0017880907,0.0004942154,0.0019125084,0.0010603504,0.0013845356,0.0022558565],"category_scores_gemma":[0.015927626,0.00033623082,0.00090884167,0.0012983308,0.0005177694,0.004604941,0.0013853203,0.002087265,0.0022552363],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001261475,0.00084795366,0.07496865,0.0005612711,0.0005720694,0.00030352373,0.0012361401,0.05924307,0.012779314,0.00559261,0.0225557,0.8200783],"study_design_scores_gemma":[0.00004426776,0.00025005286,0.011124113,0.00018154818,0.00018458009,0.00022294206,0.0005326878,0.952317,0.008331261,0.01954763,0.0071903877,0.00007349031],"about_ca_topic_score_codex":0.0033616421,"about_ca_topic_score_gemma":0.0054851994,"teacher_disagreement_score":0.0033616421,"about_ca_system_score_codex":0.00073318393,"about_ca_system_score_gemma":0.0006401083,"threshold_uncertainty_score":0.015960515},"labels":[],"label_agreement":null},{"id":"W2970259172","doi":"10.18653/v1/d19-1349","title":"Evaluating Topic Quality with Posterior Variability","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Natural language processing; Joint (building); Quality (philosophy); Artificial intelligence; Engineering; Philosophy; Epistemology","score_opus":0.08068371065089569,"score_gpt":0.36537438333173433,"score_spread":0.28469067268083864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970259172","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10737883,0.01964585,0.85277015,0.0025112526,0.000742691,0.0003239058,0.0042245286,0.0071842084,0.005218611],"genre_scores_gemma":[0.7899425,0.003452097,0.1825381,0.00047195732,0.001644753,0.00028935372,0.016377749,0.002052193,0.0032312805],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9863781,0.00688052,0.0009349474,0.0032460003,0.0020171881,0.0005433308],"domain_scores_gemma":[0.90552336,0.079359144,0.0021123984,0.0073244493,0.004195875,0.0014847937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032585762,0.0023072637,0.0031927514,0.0073799277,0.0015512021,0.0070339385,0.0024218168,0.0047788722,0.0034123072],"category_scores_gemma":[0.10463599,0.0014785948,0.0025156992,0.0042179725,0.0015639926,0.007749887,0.004241899,0.0040150997,0.0024230913],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004853325,0.00034987053,0.051626414,0.0016358405,0.0031310627,0.00038064204,0.0011862455,0.28783843,0.006811434,0.01558537,0.049983673,0.5766177],"study_design_scores_gemma":[0.00019392408,0.00017836233,0.0057971305,0.00015211929,0.0005696795,0.00023952617,0.0001713064,0.9586061,0.0029942314,0.026670437,0.004347502,0.0000795739],"about_ca_topic_score_codex":0.007013622,"about_ca_topic_score_gemma":0.005604623,"teacher_disagreement_score":0.032585762,"about_ca_system_score_codex":0.0015452346,"about_ca_system_score_gemma":0.0019607355,"threshold_uncertainty_score":0.17233199},"labels":[],"label_agreement":null},{"id":"W2970263339","doi":"10.18653/v1/d19-1298","title":"Extractive Summarization of Long Documents by Combining Global and Local Context","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":159,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Huawei Technologies","keywords":"Automatic summarization; Computer science; Context (archaeology); Information retrieval; Multi-document summarization; Geography; Archaeology","score_opus":0.008216913181033503,"score_gpt":0.2435327832202689,"score_spread":0.23531587003923538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970263339","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06253787,0.035440784,0.856269,0.0019698353,0.002266575,0.00077326875,0.016635476,0.017484333,0.006622832],"genre_scores_gemma":[0.21859758,0.009258931,0.6942616,0.00041591775,0.0029911343,0.00078564155,0.059319887,0.0015694051,0.0127998255],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988206,0.0002674211,0.00015965519,0.00034178683,0.00029359933,0.00011691433],"domain_scores_gemma":[0.9964522,0.0013067429,0.00034853572,0.00047319868,0.0012737804,0.00014559978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015951829,0.0030423831,0.0020418463,0.007648139,0.0009761564,0.0028390978,0.0014329848,0.0011865095,0.0036654375],"category_scores_gemma":[0.004388695,0.0008051693,0.0014512581,0.005899866,0.00040864456,0.003323015,0.001744017,0.0014589139,0.0053236084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008423682,0.00023076507,0.0018155321,0.0017647428,0.00043410345,0.00052857824,0.0006000651,0.0070198826,0.056047034,0.002288471,0.057456452,0.870972],"study_design_scores_gemma":[0.00055557897,0.0021043958,0.023109091,0.0010266206,0.0048084836,0.0015690319,0.0035932795,0.5074451,0.16069682,0.05717281,0.23742881,0.0004900125],"about_ca_topic_score_codex":0.0029748403,"about_ca_topic_score_gemma":0.008040829,"teacher_disagreement_score":0.007648139,"about_ca_system_score_codex":0.00048137017,"about_ca_system_score_gemma":0.0013984399,"threshold_uncertainty_score":0.012262166},"labels":[],"label_agreement":null},{"id":"W2970344992","doi":"10.2139/ssrn.3304756","title":"Focused Concept Miner (FCM): Interpretable Deep Learning for Text Exploration","year":2018,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Quest University Canada","funders":"","keywords":"Artificial intelligence; Deep learning; Computer science; Natural language processing; Psychology; Machine learning","score_opus":0.015196180519923854,"score_gpt":0.2513088160287952,"score_spread":0.23611263550887132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970344992","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013532741,0.0015207849,0.9565242,0.0005903084,0.00025148303,0.00029859197,0.0036205116,0.021761423,0.0019000288],"genre_scores_gemma":[0.13856949,0.00085109525,0.84626544,0.00047847492,0.00016978274,0.000641243,0.006876123,0.0009657213,0.00518256],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991967,0.00024923266,0.00006200414,0.00025648976,0.00017059085,0.000064944215],"domain_scores_gemma":[0.996711,0.0023443568,0.00010066186,0.00040037185,0.000313747,0.00012981916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002316111,0.0018898867,0.0010156444,0.002560413,0.0005342448,0.0014850469,0.0030111622,0.0023339076,0.010315479],"category_scores_gemma":[0.009147974,0.0009134201,0.0015681133,0.0019470665,0.0004987845,0.003355022,0.0031569882,0.003484348,0.0035603584],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034081435,0.00029664193,0.001342007,0.0006160983,0.00017625843,0.00014877267,0.0003263068,0.030276746,0.006339661,0.009266632,0.03716974,0.91370034],"study_design_scores_gemma":[0.00007482158,0.000104655795,0.0003075967,0.00010529611,0.00005710455,0.0000836146,0.000067363035,0.95099556,0.006446427,0.03200824,0.009727072,0.000022391187],"about_ca_topic_score_codex":0.003951359,"about_ca_topic_score_gemma":0.009739286,"teacher_disagreement_score":0.010315479,"about_ca_system_score_codex":0.0007926093,"about_ca_system_score_gemma":0.0018082226,"threshold_uncertainty_score":0.034508705},"labels":[],"label_agreement":null},{"id":"W2970434547","doi":"","title":"Entity and Event Extraction from Scratch Using Minimal Training Data.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Scratch; Computer science; Event (particle physics); Artificial intelligence; Extraction (chemistry); Training set; Natural language processing; Programming language; Chromatography","score_opus":0.050751457151561076,"score_gpt":0.3190120562127388,"score_spread":0.2682605990611777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970434547","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03557865,0.0019673277,0.9012777,0.00081029313,0.00041035042,0.0007429458,0.03518319,0.015242102,0.008787399],"genre_scores_gemma":[0.22084652,0.0013414988,0.58512574,0.0003454944,0.00029110943,0.0011172533,0.18229158,0.0008279941,0.007812914],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99803096,0.00039190735,0.00019724801,0.0009028395,0.00029549532,0.00018148324],"domain_scores_gemma":[0.99349505,0.0033035853,0.00024729394,0.0018934131,0.0008715487,0.00018915226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001762571,0.0016132524,0.0012396224,0.0042817546,0.0012378731,0.0021954286,0.002609085,0.001818047,0.0080146585],"category_scores_gemma":[0.012988849,0.0007680905,0.0018142099,0.0043970826,0.0005659687,0.0057704444,0.0028000122,0.00235238,0.011510357],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007009722,0.0006634981,0.010638247,0.0016458464,0.00039775565,0.00089633145,0.0005788075,0.017228222,0.025280213,0.011038722,0.09097542,0.8399559],"study_design_scores_gemma":[0.00018337881,0.0004672464,0.024950705,0.00084953767,0.0009099562,0.0025044233,0.0018887583,0.59461427,0.05742871,0.10864832,0.2073611,0.00019356105],"about_ca_topic_score_codex":0.0058768354,"about_ca_topic_score_gemma":0.013058691,"teacher_disagreement_score":0.0080146585,"about_ca_system_score_codex":0.0006729538,"about_ca_system_score_gemma":0.0027262205,"threshold_uncertainty_score":0.02681166},"labels":[],"label_agreement":null},{"id":"W2970435863","doi":"","title":"GAIA - A Multi-media Multi-lingual Knowledge Extraction and Hypothesis Generation System.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Extraction (chemistry); Artificial intelligence; Natural language processing; Chromatography; Chemistry","score_opus":0.04184139271989265,"score_gpt":0.2906546889050337,"score_spread":0.24881329618514106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970435863","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027391309,0.0029968354,0.7291967,0.0017258617,0.0008735718,0.0020671606,0.045810774,0.17976697,0.010170769],"genre_scores_gemma":[0.09602879,0.0007279186,0.83312637,0.0008325258,0.00028600259,0.001995653,0.059040044,0.0018672161,0.006095488],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99765766,0.0009423448,0.0002094398,0.0006980016,0.00037414994,0.00011836261],"domain_scores_gemma":[0.99272376,0.004851185,0.00033955433,0.0007782512,0.00097308774,0.00033413162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00281606,0.0018524448,0.0014007618,0.0066773486,0.0011239392,0.0028360914,0.0026056352,0.0022221506,0.010194306],"category_scores_gemma":[0.013943223,0.00072547636,0.0017468185,0.0032656624,0.0006194867,0.005585749,0.0042791595,0.0022037223,0.012508686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011305735,0.0006641188,0.005766254,0.002412366,0.0010904394,0.0014857656,0.001051915,0.0065671755,0.031688284,0.007119084,0.19320811,0.747816],"study_design_scores_gemma":[0.0009187351,0.00091496576,0.016468465,0.0006389163,0.0016978579,0.0020863744,0.002637598,0.56865704,0.06427513,0.08233306,0.25883928,0.00053263566],"about_ca_topic_score_codex":0.002833106,"about_ca_topic_score_gemma":0.00477493,"teacher_disagreement_score":0.010194306,"about_ca_system_score_codex":0.0006519606,"about_ca_system_score_gemma":0.001966466,"threshold_uncertainty_score":0.034103334},"labels":[],"label_agreement":null},{"id":"W2970600560","doi":"10.18653/v1/d19-1620","title":"Countering the Effects of Lead Bias in News Summarization via Multi-Stage Training and Auxiliary Losses","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Institut de Valorisation des Données; Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Automatic summarization; Computer science; Natural language processing; Artificial intelligence; Natural (archaeology); Joint (building); Natural language; Speech recognition; Engineering; History","score_opus":0.06762744025516274,"score_gpt":0.28303542772080464,"score_spread":0.2154079874656419,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970600560","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1301457,0.0039350707,0.8518272,0.0028942046,0.00053941837,0.00017664074,0.0005452733,0.0048198,0.005116646],"genre_scores_gemma":[0.7979117,0.00070669444,0.18682113,0.0010571408,0.00091894035,0.00026301705,0.0017224579,0.0009962282,0.009602643],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99687237,0.0019419267,0.00016402682,0.000522501,0.00028680867,0.00021242503],"domain_scores_gemma":[0.9669429,0.02678842,0.0010527382,0.0023318992,0.0023394192,0.00054457557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010741287,0.0016062617,0.001768504,0.0012410051,0.00096317613,0.0016127272,0.0022716096,0.002443939,0.0034758598],"category_scores_gemma":[0.040093504,0.000997188,0.0007396336,0.0010056156,0.0010320152,0.00549692,0.0032690696,0.0041251415,0.0019026824],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004509738,0.00073765416,0.0066352077,0.00071697764,0.00048628057,0.0002808685,0.0010367862,0.24633074,0.02110061,0.021752542,0.027357884,0.6690547],"study_design_scores_gemma":[0.00019374036,0.00034162798,0.0012128518,0.000048695427,0.0001242507,0.00005436939,0.000071059716,0.973328,0.006347025,0.01680147,0.001450497,0.00002635091],"about_ca_topic_score_codex":0.0026488702,"about_ca_topic_score_gemma":0.0058408193,"teacher_disagreement_score":0.010741287,"about_ca_system_score_codex":0.0007924628,"about_ca_system_score_gemma":0.0013360414,"threshold_uncertainty_score":0.056806087},"labels":[],"label_agreement":null},{"id":"W2970696237","doi":"","title":"The UTexas system for TAC SM-KBP task 3: Probabilistic generation of coherent hypotheses.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Probabilistic logic; Task (project management); Computer science; Artificial intelligence; Systems engineering; Engineering","score_opus":0.026984973425116365,"score_gpt":0.2557323405725275,"score_spread":0.22874736714741112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970696237","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044586558,0.0017339,0.56531256,0.0016921585,0.00081717677,0.0017893274,0.07569532,0.2839206,0.024452332],"genre_scores_gemma":[0.23010153,0.00049027003,0.5729899,0.00065505167,0.00046335993,0.002269321,0.17144409,0.008847328,0.0127392225],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99815065,0.0007574048,0.00014735285,0.00051668746,0.0003316271,0.000096358985],"domain_scores_gemma":[0.99470407,0.003422438,0.00020057596,0.00074095535,0.00069932017,0.00023247612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025607357,0.0020404516,0.0013469344,0.002203542,0.0010063013,0.0025488387,0.0027633733,0.0023184216,0.03606809],"category_scores_gemma":[0.017928498,0.00066314824,0.0011807143,0.0012699643,0.00048322845,0.0044637495,0.0037635306,0.0017861675,0.018675748],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019764125,0.00035358075,0.0035771383,0.002697445,0.0004502371,0.000978034,0.0017848298,0.012309324,0.017441396,0.011943729,0.5325755,0.41391242],"study_design_scores_gemma":[0.0009521558,0.0004952465,0.0044700443,0.00045665726,0.00036507414,0.001300852,0.0016906234,0.6900912,0.03797977,0.0518303,0.21015403,0.00021400984],"about_ca_topic_score_codex":0.006051705,"about_ca_topic_score_gemma":0.008416992,"teacher_disagreement_score":0.03606809,"about_ca_system_score_codex":0.0007908183,"about_ca_system_score_gemma":0.0020152829,"threshold_uncertainty_score":0.12065977},"labels":[],"label_agreement":null},{"id":"W2970808735","doi":"10.18653/v1/d19-1069","title":"KnowledgeNet: A Benchmark Dataset for Knowledge Base Population","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Benchmark (surveying); Computer science; Natural language processing; Base (topology); Population; Artificial intelligence; Knowledge base; Joint (building); Geography; Cartography; Engineering; Demography; Mathematics; Sociology","score_opus":0.03238429445618977,"score_gpt":0.28857149760848305,"score_spread":0.25618720315229326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970808735","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02616194,0.0047074147,0.012094767,0.0022116401,0.00059035566,0.00067996193,0.92216575,0.015158198,0.016229898],"genre_scores_gemma":[0.014797847,0.0008629343,0.017045615,0.00026675616,0.000044501565,0.00047296754,0.96393657,0.00030041207,0.0022723519],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974782,0.0005104754,0.00039538555,0.00064189034,0.0007930235,0.00018107801],"domain_scores_gemma":[0.9933954,0.0027136377,0.00038535052,0.0013039939,0.0016309993,0.00057058374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031276231,0.0023702325,0.0012353507,0.009242417,0.0016329739,0.0030454018,0.0047727255,0.003991804,0.010137223],"category_scores_gemma":[0.018311178,0.0006811331,0.0015705341,0.008099057,0.0007078666,0.004478029,0.002462706,0.0021052293,0.010313682],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044791907,0.00055722264,0.004978555,0.0024438892,0.0003023562,0.00036938666,0.00018326023,0.011477937,0.0009982014,0.005501972,0.8925943,0.08014516],"study_design_scores_gemma":[0.0010480256,0.0004001786,0.014163128,0.0013178932,0.0004641108,0.0010039671,0.00083111314,0.089620024,0.007894871,0.021282198,0.8617993,0.00017527031],"about_ca_topic_score_codex":0.037333433,"about_ca_topic_score_gemma":0.04731642,"teacher_disagreement_score":0.037333433,"about_ca_system_score_codex":0.002772368,"about_ca_system_score_gemma":0.004647187,"threshold_uncertainty_score":0.07423222},"labels":[],"label_agreement":null},{"id":"W2970844223","doi":"","title":"Team EP at TAC 2018: Automating data extraction in systematic reviews of environmental agents.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Mistake; Computer science; Regularization (linguistics); Task (project management); Word (group theory); Artificial intelligence; Natural language processing; Set (abstract data type); Training set; Data extraction; Sequence (biology); Layer (electronics); Information retrieval; Machine learning; Programming language","score_opus":0.039178447338124624,"score_gpt":0.30054230997182163,"score_spread":0.261363862633697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970844223","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012900534,0.028362801,0.32323998,0.012395164,0.0046336693,0.017370263,0.46345142,0.11683531,0.020810807],"genre_scores_gemma":[0.025103856,0.0041687093,0.73502433,0.002378224,0.00062243367,0.01334104,0.20434529,0.0047076414,0.0103084585],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96843976,0.016640153,0.0054608537,0.0033720292,0.0055434257,0.0005437753],"domain_scores_gemma":[0.8704247,0.0759238,0.014363349,0.015423967,0.0202598,0.0036044717],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.046684813,0.0022759433,0.0024173593,0.014804368,0.0018387581,0.0045924704,0.002719084,0.0022020433,0.03215055],"category_scores_gemma":[0.17240511,0.0014962329,0.003620861,0.007890671,0.0004387391,0.005229298,0.006389199,0.0022521536,0.023264669],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004893586,0.0001240922,0.0037846074,0.034649048,0.0014040337,0.00028721156,0.0012231014,0.0010971446,0.0054006698,0.002299,0.6311459,0.31809583],"study_design_scores_gemma":[0.00095529633,0.00061799504,0.013836797,0.013375294,0.0019363468,0.0006251334,0.0011355899,0.017518988,0.01515433,0.016107904,0.91840756,0.00032886546],"about_ca_topic_score_codex":0.006031378,"about_ca_topic_score_gemma":0.026491577,"teacher_disagreement_score":0.9533152,"about_ca_system_score_codex":0.0024786075,"about_ca_system_score_gemma":0.0151080405,"threshold_uncertainty_score":0.24689585},"labels":[],"label_agreement":null},{"id":"W2970877041","doi":"","title":"Overview of the TAC 2018 Systematic Review Information Extraction Track.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Track (disk drive); Computer science; Extraction (chemistry); Information extraction; Information retrieval; Chromatography; Operating system; Chemistry","score_opus":0.019582331816022865,"score_gpt":0.28274399945225454,"score_spread":0.26316166763623167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970877041","genre_codex":"review","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041637244,0.4335758,0.07297111,0.02122517,0.004578206,0.022763222,0.40640327,0.010811261,0.023508156],"genre_scores_gemma":[0.016000528,0.2651452,0.18032229,0.0132309245,0.002756712,0.05403164,0.45038223,0.0028426994,0.015287821],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9696218,0.0082609365,0.011929905,0.00340687,0.0061219498,0.0006584551],"domain_scores_gemma":[0.8106486,0.08712403,0.03541655,0.016778318,0.045478042,0.004554402],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05587626,0.002044382,0.0065861675,0.0589611,0.0029336195,0.009512066,0.00555538,0.0028029215,0.0355476],"category_scores_gemma":[0.18609771,0.0017295033,0.0065993434,0.03921993,0.001166549,0.0074515888,0.007069309,0.0025766506,0.017732153],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011322482,0.00010121237,0.002456758,0.31426737,0.004447943,0.00018891349,0.0013907594,0.00085086067,0.0028893994,0.005239244,0.36104086,0.30599445],"study_design_scores_gemma":[0.0005068904,0.00035040372,0.005474403,0.16614468,0.008714843,0.00027626505,0.00033810936,0.0008755912,0.0018576696,0.007997448,0.8072502,0.0002135039],"about_ca_topic_score_codex":0.012622684,"about_ca_topic_score_gemma":0.040990844,"teacher_disagreement_score":0.0589611,"about_ca_system_score_codex":0.0064411894,"about_ca_system_score_gemma":0.042297017,"threshold_uncertainty_score":0.2955054},"labels":[],"label_agreement":null},{"id":"W2971001654","doi":"10.18653/v1/d19-1451","title":"Aligning Cross-Lingual Entities with Multi-Aspect Information","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":167,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Ministry of Education, India; Singapore University of Technology and Design; Compute Canada","keywords":"Computer science; Natural language processing; Joint (building); Natural (archaeology); Artificial intelligence; Engineering; History; Archaeology; Architectural engineering","score_opus":0.013448448807038492,"score_gpt":0.25059866825851046,"score_spread":0.23715021945147197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971001654","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16555026,0.008350167,0.76541686,0.0014703725,0.0018135338,0.00037455125,0.022066338,0.015783606,0.019174293],"genre_scores_gemma":[0.55276394,0.0023257814,0.3663178,0.0005316083,0.0004492012,0.00033613213,0.06630371,0.0034374169,0.007534437],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997914,0.00048749946,0.00018610833,0.00085892703,0.0003548681,0.00019868743],"domain_scores_gemma":[0.99720997,0.00070150394,0.00027560876,0.00086112635,0.0008114382,0.00014039996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016731062,0.0012585195,0.0012368367,0.00477714,0.0014012114,0.002743062,0.001462424,0.00134427,0.0039308527],"category_scores_gemma":[0.0052754,0.0008690088,0.0017610346,0.008064916,0.00055018143,0.006720831,0.0040432992,0.0017747747,0.0045154244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012122415,0.00040041795,0.03804678,0.0013551392,0.0010526433,0.0018931603,0.0021411777,0.013646461,0.058963332,0.021710139,0.10104899,0.7585295],"study_design_scores_gemma":[0.00019525163,0.00033718866,0.0475247,0.00040959945,0.0019095696,0.0020786677,0.0044858567,0.48875245,0.06777406,0.08483774,0.3013314,0.0003635043],"about_ca_topic_score_codex":0.008212684,"about_ca_topic_score_gemma":0.02123898,"teacher_disagreement_score":0.008212684,"about_ca_system_score_codex":0.0006147704,"about_ca_system_score_gemma":0.0015376864,"threshold_uncertainty_score":0.016329765},"labels":[],"label_agreement":null},{"id":"W2971048662","doi":"10.18653/v1/d19-1403","title":"Induction Networks for Few-Shot Text Classification","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":201,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Computer science; Ping (video games); Natural language processing; Shot (pellet); Artificial intelligence; Joint (building); Computer security; Engineering; Chemistry","score_opus":0.07470951185655068,"score_gpt":0.2798845743633619,"score_spread":0.20517506250681122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971048662","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021109233,0.0046401564,0.95585036,0.0010369676,0.0005380843,0.00023639342,0.0028122317,0.010153463,0.0036231144],"genre_scores_gemma":[0.5254271,0.0028005983,0.4142004,0.0010555346,0.0013391976,0.0011819449,0.023974705,0.0012558796,0.028764537],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99851185,0.0003506923,0.000103890095,0.00059672,0.00025010723,0.00018672469],"domain_scores_gemma":[0.99647516,0.0020196198,0.00024571246,0.0005298071,0.00057692477,0.00015278764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015592277,0.001653233,0.00223042,0.0036530609,0.0015974074,0.001483748,0.003738035,0.0025673965,0.006905869],"category_scores_gemma":[0.006119912,0.0010110578,0.0017421709,0.0027949414,0.0007458083,0.004159567,0.0020761774,0.003441118,0.005772803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058998645,0.00042307156,0.002955102,0.0005088663,0.000314734,0.00030217605,0.00021656865,0.13643491,0.0055491505,0.014326168,0.04124043,0.7971388],"study_design_scores_gemma":[0.000011327247,0.000027705388,0.00034927047,0.000020924535,0.000034608063,0.000042523567,0.000021322325,0.973224,0.0011509237,0.023126567,0.0019794665,0.0000114273835],"about_ca_topic_score_codex":0.009542249,"about_ca_topic_score_gemma":0.014873189,"teacher_disagreement_score":0.009542249,"about_ca_system_score_codex":0.00157783,"about_ca_system_score_gemma":0.0012557864,"threshold_uncertainty_score":0.023102462},"labels":[],"label_agreement":null},{"id":"W2971160449","doi":"10.18653/v1/d19-1272","title":"Keep Calm and Switch On! Preserving Sentiment and Fluency in Semantic Text Exchange","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Huawei Technologies","keywords":"Fluency; Natural (archaeology); Natural language; Semantics (computer science); Joint (building)","score_opus":0.013427869232149468,"score_gpt":0.2322421173530616,"score_spread":0.21881424812091213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971160449","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46976328,0.0021873014,0.46787393,0.006965866,0.001395425,0.00022840584,0.0019368847,0.00391755,0.045731343],"genre_scores_gemma":[0.9498586,0.00025614473,0.042271417,0.0004583726,0.0003626548,0.00006735582,0.0008726625,0.00048793646,0.0053648297],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985343,0.00076560595,0.00006904234,0.00031962295,0.00016058674,0.00015078997],"domain_scores_gemma":[0.99449956,0.0025093188,0.00047198462,0.0014276317,0.0007744671,0.00031703414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024217567,0.0004922387,0.00058140105,0.0006799984,0.0012646412,0.002513501,0.00067714514,0.0010586744,0.0052455952],"category_scores_gemma":[0.017016996,0.0004201232,0.0004767455,0.0007234907,0.0014192112,0.009604961,0.0023750858,0.001214192,0.0029068207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0045650136,0.0005094576,0.018548457,0.0008080959,0.0003438378,0.0006884568,0.0071115736,0.0145661095,0.048791625,0.16344318,0.051517643,0.6891066],"study_design_scores_gemma":[0.00025077155,0.00052020676,0.018210137,0.00015633342,0.00034900688,0.00075128756,0.004114411,0.17392457,0.023829507,0.7378193,0.0398531,0.00022135746],"about_ca_topic_score_codex":0.0009498165,"about_ca_topic_score_gemma":0.0014090181,"teacher_disagreement_score":0.0052455952,"about_ca_system_score_codex":0.00043900288,"about_ca_system_score_gemma":0.00037053155,"threshold_uncertainty_score":0.017548263},"labels":[],"label_agreement":null},{"id":"W2971569798","doi":"10.18653/v1/d19-1006","title":"How Contextual are Contextualized Word Representations? Comparing the Geometry of BERT, ELMo, and GPT-2 Embeddings","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Word (group theory); Similarity (geometry); Context (archaeology); Computer science; Embedding; Cosine similarity; Natural language processing; Task (project management); Artificial intelligence; Variance (accounting); Linguistics; Pattern recognition (psychology); Image (mathematics); History; Philosophy","score_opus":0.05486834225772868,"score_gpt":0.29590379590506066,"score_spread":0.24103545364733198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971569798","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2592553,0.0019761755,0.72072285,0.0027018443,0.00017506156,0.00007480948,0.0013003598,0.0017401052,0.012053493],"genre_scores_gemma":[0.88407123,0.0010383032,0.10780292,0.0003387462,0.000104905266,0.00014292226,0.0019118392,0.00042059907,0.0041686073],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99958235,0.00013931542,0.000020597781,0.000134099,0.000060271454,0.00006329345],"domain_scores_gemma":[0.99880064,0.00051856384,0.00010403794,0.00028508532,0.00019679565,0.000094849755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007347153,0.00079114275,0.0005750758,0.00080346706,0.0004279216,0.0015120014,0.0008161032,0.0010629168,0.0036668605],"category_scores_gemma":[0.0070296666,0.00055857364,0.0007164,0.0007825188,0.0012365456,0.005083159,0.0016957095,0.0018488307,0.00090052],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070667814,0.0000840931,0.0072294446,0.00041411625,0.00022044554,0.00037015896,0.0014131796,0.36265194,0.018485675,0.29558048,0.00998347,0.30286035],"study_design_scores_gemma":[0.000038588194,0.000110764355,0.002528853,0.000056999746,0.00005662353,0.00020625828,0.0002729897,0.719627,0.0037394578,0.26800427,0.0053141643,0.000043985125],"about_ca_topic_score_codex":0.00530776,"about_ca_topic_score_gemma":0.007069892,"teacher_disagreement_score":0.00530776,"about_ca_system_score_codex":0.00095881836,"about_ca_system_score_gemma":0.00069417217,"threshold_uncertainty_score":0.012266815},"labels":[],"label_agreement":null},{"id":"W2972479880","doi":"10.18653/v1/w19-3221","title":"Detection of Adverse Drug Reaction Mentions in Tweets Using ELMo","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Public Health","funders":"","keywords":"Computer science; Natural language processing; Task (project management); Lexicon; Similarity (geometry); Representation (politics); Artificial intelligence; Matching (statistics); Word (group theory); Information retrieval; Adverse drug reaction; Semantic similarity; Drug; Linguistics; Medicine; Image (mathematics)","score_opus":0.02021065579365378,"score_gpt":0.25405002550630984,"score_spread":0.23383936971265606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972479880","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49323004,0.007160143,0.3177297,0.0037427694,0.0022035963,0.001789783,0.10500679,0.050757766,0.018379342],"genre_scores_gemma":[0.68057525,0.0011465085,0.18949996,0.0010623308,0.00054709485,0.00095596834,0.11136867,0.00076921156,0.014074992],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990621,0.00020590842,0.00010388521,0.00034390457,0.00016509695,0.00011921066],"domain_scores_gemma":[0.9986546,0.0006242275,0.00014077376,0.00017863518,0.00033697163,0.000064771084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015263116,0.0019991696,0.00085197826,0.003931581,0.00055941707,0.001209141,0.0009616265,0.0015314886,0.003825127],"category_scores_gemma":[0.0040340465,0.0004149333,0.0018154451,0.0012996044,0.00024671367,0.001625713,0.0016637786,0.001271694,0.004309414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024499693,0.0015958142,0.080769114,0.0021799156,0.0013464311,0.0016938676,0.0008502377,0.064529985,0.043809436,0.004878109,0.11070119,0.685196],"study_design_scores_gemma":[0.000105949934,0.00047106258,0.02433846,0.00015689565,0.0003867892,0.00069294637,0.0003227653,0.90303206,0.022471866,0.0069699655,0.04094816,0.000103049446],"about_ca_topic_score_codex":0.0061826976,"about_ca_topic_score_gemma":0.011015223,"teacher_disagreement_score":0.0061826976,"about_ca_system_score_codex":0.0006996328,"about_ca_system_score_gemma":0.0009014498,"threshold_uncertainty_score":0.012796342},"labels":[],"label_agreement":null},{"id":"W2972603547","doi":"10.18653/v1/w19-4115","title":"Relevant and Informative Response Generation using Pointwise Mutual Information","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Japan Society for the Promotion of Science; Microsoft Research Asia; Microsoft Research","keywords":"Pointwise; Pointwise mutual information; Computer science; Utterance; Sequence (biology); Mutual information; Simple (philosophy); Artificial intelligence; Machine learning; Mathematics","score_opus":0.021586017779013967,"score_gpt":0.23908199842080233,"score_spread":0.21749598064178835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972603547","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06295878,0.00022058547,0.93129873,0.0004472435,0.000050804083,0.00018888719,0.00011782579,0.0017233108,0.0029939201],"genre_scores_gemma":[0.8410605,0.00013780482,0.15395938,0.0002448175,0.00006101776,0.0004215919,0.00027137026,0.00026267662,0.0035807926],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974069,0.0015703955,0.0000850908,0.00048491283,0.0003319072,0.00012077556],"domain_scores_gemma":[0.9935214,0.0048392266,0.00034381886,0.0004897292,0.00061065896,0.00019521704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029106017,0.0010571525,0.0007656092,0.00084890425,0.00039917638,0.0007786731,0.0012699413,0.0013386792,0.0028337366],"category_scores_gemma":[0.011889505,0.00048169104,0.00072786066,0.00045913388,0.0008828667,0.0019582915,0.0013532109,0.0014236396,0.0009885314],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012289436,0.000498433,0.0043920013,0.00045243715,0.00025119892,0.00041335064,0.0014162879,0.49023575,0.054165907,0.034730252,0.0059586535,0.40625682],"study_design_scores_gemma":[0.00002608892,0.000096906304,0.00029071575,0.00001089109,0.000021510277,0.00006819648,0.000027950513,0.98231137,0.005803017,0.010746497,0.0005768476,0.000020042195],"about_ca_topic_score_codex":0.0012403398,"about_ca_topic_score_gemma":0.0018613215,"teacher_disagreement_score":0.0029106017,"about_ca_system_score_codex":0.0006643447,"about_ca_system_score_gemma":0.0008764919,"threshold_uncertainty_score":0.0153928995},"labels":[],"label_agreement":null},{"id":"W2972703796","doi":"10.48550/arxiv.1909.05246","title":"Self-Attentional Models Application in Task-Oriented Dialogue Generation Systems","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Task (project management); Sequence (biology); Machine translation; Artificial intelligence; Mechanism (biology); Sequence learning; Convolution (computer science); Machine learning; Natural language processing; Artificial neural network; Engineering","score_opus":0.06918778667700219,"score_gpt":0.18057491312508084,"score_spread":0.11138712644807865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972703796","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1767988,0.001684359,0.8022703,0.000695101,0.0002948294,0.00028713714,0.0005535415,0.011214329,0.0062015923],"genre_scores_gemma":[0.82993674,0.00031716144,0.16354649,0.0003150506,0.00009609949,0.00032147253,0.0011358539,0.0003960303,0.0039351517],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983266,0.0008562391,0.00009398966,0.0004157102,0.00020567678,0.00010181351],"domain_scores_gemma":[0.9966633,0.002070104,0.00016591584,0.00039309845,0.00054772315,0.00015995631],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031439266,0.0008784366,0.0005793013,0.0008435199,0.00045158435,0.0011976149,0.0010539371,0.0014092488,0.002237228],"category_scores_gemma":[0.007609804,0.00035759958,0.0008282168,0.00049077574,0.0004043155,0.001734578,0.0013522814,0.0014113622,0.0010577479],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009065685,0.0006917777,0.005861432,0.0007001755,0.00037805905,0.00040026673,0.0021940246,0.43646595,0.041948866,0.011174691,0.01129756,0.48798066],"study_design_scores_gemma":[0.00002197085,0.00010988024,0.00056320976,0.000016251757,0.00003256745,0.000049675342,0.000057348316,0.9869876,0.0066472767,0.0034398737,0.0020576068,0.000016763102],"about_ca_topic_score_codex":0.0046063266,"about_ca_topic_score_gemma":0.005441187,"teacher_disagreement_score":0.0046063266,"about_ca_system_score_codex":0.00074047956,"about_ca_system_score_gemma":0.0008424645,"threshold_uncertainty_score":0.016626835},"labels":[],"label_agreement":null},{"id":"W2972715831","doi":"10.1016/j.neucom.2019.08.082","title":"Finding decision jumps in text classification","year":2019,"lang":"en","type":"article","venue":"Neurocomputing","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Natural Science Foundation of China","keywords":"Computer science; Artificial intelligence; Jumper; Machine learning; Benchmark (surveying); Process (computing); Artificial neural network; Feature (linguistics); Reinforcement learning; Key (lock); Reading (process)","score_opus":0.03601712965880887,"score_gpt":0.27722046955661817,"score_spread":0.2412033398978093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972715831","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55965924,0.0052700914,0.4234342,0.0031566268,0.00048623842,0.00029300337,0.0018434017,0.0027155334,0.0031417427],"genre_scores_gemma":[0.95255184,0.00048462715,0.04067212,0.00023249649,0.00039380448,0.0000983704,0.0024402384,0.00017857211,0.0029479072],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982482,0.00046913297,0.0001745316,0.00047450207,0.00036706668,0.00026653986],"domain_scores_gemma":[0.9699471,0.026965534,0.00084144174,0.0006384813,0.0008995159,0.0007079948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00342445,0.0008795112,0.001956433,0.0047075325,0.0014980193,0.002358176,0.0019751817,0.0033659472,0.004770138],"category_scores_gemma":[0.018164353,0.00076529995,0.0014172756,0.0027593158,0.0011265229,0.0042889286,0.0018510469,0.0040761344,0.0011325312],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0052250642,0.0015870043,0.055402633,0.0017834246,0.00060597085,0.0014236929,0.001537337,0.1793878,0.016097177,0.03860645,0.03539488,0.66294855],"study_design_scores_gemma":[0.00009442708,0.00017087288,0.003962005,0.00008206143,0.00009608145,0.00011664903,0.00023137096,0.94093597,0.0034997032,0.049500283,0.001279423,0.00003106039],"about_ca_topic_score_codex":0.0028279193,"about_ca_topic_score_gemma":0.0029322503,"teacher_disagreement_score":0.004770138,"about_ca_system_score_codex":0.0010096403,"about_ca_system_score_gemma":0.0009649906,"threshold_uncertainty_score":0.018110454},"labels":[],"label_agreement":null},{"id":"W2973276266","doi":"10.48550/arxiv.1909.09268","title":"Towards Neural Language Evaluators","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Computer science; Metric (unit); Natural language processing; Transformer; Artificial intelligence; BLEU; Machine learning; Information retrieval; Machine translation; Engineering","score_opus":0.08674456871558647,"score_gpt":0.211965579487791,"score_spread":0.12522101077220452,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973276266","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024760012,0.013118546,0.92767346,0.0069645634,0.0006648533,0.00023674828,0.0014867728,0.00953584,0.015559209],"genre_scores_gemma":[0.48135087,0.0041334485,0.48556313,0.0034531676,0.0013758262,0.0008295402,0.0049462877,0.0014227577,0.016924938],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9852775,0.008820378,0.0007280151,0.0022337202,0.0024792685,0.00046126114],"domain_scores_gemma":[0.95587045,0.02753542,0.0017128815,0.0038025323,0.009903014,0.0011756782],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019377932,0.0016947445,0.0016271145,0.0041736695,0.0009572037,0.005538515,0.0029247038,0.0024403285,0.006343194],"category_scores_gemma":[0.086011276,0.00060348643,0.00076549547,0.0022801254,0.0015701439,0.010675587,0.0035635976,0.0050997604,0.0052761585],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005040371,0.00025757987,0.005777334,0.0009250586,0.0002633281,0.0000935048,0.00075343344,0.04450248,0.0063769417,0.049254775,0.04047156,0.85081995],"study_design_scores_gemma":[0.00007860359,0.00040084982,0.0021807393,0.0006861507,0.00016736636,0.00016829983,0.00049907126,0.7662636,0.013727936,0.17787337,0.03784194,0.00011199558],"about_ca_topic_score_codex":0.0028860702,"about_ca_topic_score_gemma":0.0047917278,"teacher_disagreement_score":0.019377932,"about_ca_system_score_codex":0.0027047987,"about_ca_system_score_gemma":0.0021114273,"threshold_uncertainty_score":0.102481544},"labels":[],"label_agreement":null},{"id":"W2973459423","doi":"10.48550/arxiv.1909.08187","title":"Learning to Generate Questions with Adaptive Copying Neural Networks","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Copying; Computer science; Artificial intelligence; Component (thermodynamics); Artificial neural network; Recurrent neural network; Natural language processing; Natural language generation; Machine learning; Natural language","score_opus":0.06812994601911321,"score_gpt":0.18441406856087042,"score_spread":0.11628412254175721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973459423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08981547,0.00093367783,0.9008445,0.0005430233,0.00012639361,0.00021750148,0.00035002004,0.00461214,0.002557314],"genre_scores_gemma":[0.7841588,0.00042790978,0.20670779,0.00039538648,0.00013905106,0.00038799885,0.0013152057,0.00029947935,0.0061682356],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9992009,0.00032052028,0.00004592561,0.00026186986,0.00011650598,0.000054285305],"domain_scores_gemma":[0.99731207,0.0017835565,0.00020401838,0.00031249825,0.00030817962,0.00007974926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015049658,0.0008186646,0.000586658,0.0006562615,0.00027370185,0.00072170445,0.0016266351,0.0013368954,0.002268958],"category_scores_gemma":[0.0068222685,0.0003694334,0.000873294,0.0005318431,0.0006112794,0.0025057648,0.0012281073,0.0013932958,0.0009143789],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005965809,0.00028532962,0.0037265033,0.00040192634,0.00017488888,0.00041638015,0.0007441515,0.22885802,0.057427853,0.0138032595,0.00784002,0.6857251],"study_design_scores_gemma":[0.000023504603,0.00008524344,0.00047645922,0.000010008614,0.000027988695,0.00007207144,0.00003082891,0.9776977,0.008796069,0.01134043,0.0014249596,0.000014800662],"about_ca_topic_score_codex":0.0018293461,"about_ca_topic_score_gemma":0.0024939491,"teacher_disagreement_score":0.002268958,"about_ca_system_score_codex":0.00063635065,"about_ca_system_score_gemma":0.0004243082,"threshold_uncertainty_score":0.007959127},"labels":[],"label_agreement":null},{"id":"W2973475963","doi":"10.18653/v1/d19-1389","title":"BottleSum: Unsupervised and Self-supervised Sentence Summarization using the Information Bottleneck Principle","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Army Research Office; Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; Samsung; Allen Institute for Artificial Intelligence; Naval Information Warfare Center Pacific; National Science Foundation","keywords":"Automatic summarization; Computer science; Sentence; Bottleneck; Artificial intelligence; Natural language processing; Information bottleneck method; Language model; Machine learning; Cluster analysis","score_opus":0.027054068520040685,"score_gpt":0.2545242226792956,"score_spread":0.2274701541592549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973475963","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007117646,0.0013951085,0.96685725,0.00036660364,0.0004817496,0.0003167301,0.0026070336,0.019311758,0.0015461461],"genre_scores_gemma":[0.07677729,0.0006224163,0.8846545,0.00035208152,0.00063039793,0.00087802566,0.021902137,0.003490126,0.01069304],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997867,0.0008240759,0.0001423564,0.00057168445,0.00045540227,0.0001394272],"domain_scores_gemma":[0.99632853,0.0017282874,0.0002113109,0.00068056956,0.00090704055,0.00014420049],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035509043,0.0024902285,0.0020779567,0.0031545372,0.0011146314,0.0021825284,0.0038043712,0.0017407973,0.0071389084],"category_scores_gemma":[0.0079790745,0.00092647306,0.0018673802,0.002297244,0.0006502266,0.0040641907,0.0031406826,0.002122241,0.006764681],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009379241,0.0003707653,0.00062613474,0.0008078774,0.0005643051,0.00022150076,0.00041855988,0.03162588,0.021696681,0.0076751644,0.09799638,0.83705884],"study_design_scores_gemma":[0.0002731114,0.00033964278,0.0011350281,0.000080497935,0.000270032,0.00013600664,0.00020900572,0.9054461,0.025385784,0.037268702,0.02934709,0.000108893975],"about_ca_topic_score_codex":0.0032480177,"about_ca_topic_score_gemma":0.006341075,"teacher_disagreement_score":0.0071389084,"about_ca_system_score_codex":0.0006545219,"about_ca_system_score_gemma":0.0016367713,"threshold_uncertainty_score":0.023882031},"labels":[],"label_agreement":null},{"id":"W2973837416","doi":"10.1145/3342558.3345404","title":"Impact of In-domain Vector Representations on the Classification of Disease-related Tweets","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Computer science; Word embedding; Artificial intelligence; Sentiment analysis; Word (group theory); Natural language processing; Domain (mathematical analysis); Task (project management); Convolutional neural network; Embedding; Initialization; Machine learning; Mathematics","score_opus":0.0348331680969018,"score_gpt":0.305411897964036,"score_spread":0.2705787298671342,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973837416","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88212806,0.002890209,0.106991075,0.0009846983,0.00034702508,0.00017080805,0.000966903,0.0014935707,0.0040276116],"genre_scores_gemma":[0.9701203,0.00055244553,0.026934912,0.00008767824,0.000060222977,0.00004691002,0.0013208477,0.00004878168,0.000827835],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975852,0.0013685089,0.00022065897,0.00036005068,0.00027773212,0.00018786023],"domain_scores_gemma":[0.9905129,0.007111822,0.0005176643,0.0006286599,0.0010339221,0.00019509683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048035528,0.0015452087,0.0008802346,0.0013289156,0.00039893348,0.0017326757,0.0005926309,0.0010503792,0.0010690449],"category_scores_gemma":[0.016946184,0.00025108838,0.0006055849,0.0013634118,0.000476574,0.0029893834,0.001130849,0.0014910571,0.000579446],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002618794,0.0017324537,0.04365767,0.00051245413,0.00056193984,0.0002406802,0.00037250057,0.21950562,0.022760544,0.0025269117,0.006730913,0.6987795],"study_design_scores_gemma":[0.00003710021,0.00034526023,0.0052074175,0.000034759447,0.000107786866,0.000095784,0.00025086256,0.9836971,0.007920677,0.0015450213,0.00073286536,0.000025417658],"about_ca_topic_score_codex":0.004025322,"about_ca_topic_score_gemma":0.0033765337,"teacher_disagreement_score":0.0048035528,"about_ca_system_score_codex":0.0006343467,"about_ca_system_score_gemma":0.00067974307,"threshold_uncertainty_score":0.025403917},"labels":[],"label_agreement":null},{"id":"W2973947483","doi":"10.1145/3345557","title":"Question Answering in Knowledge Bases","year":2019,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Dream Project of Ministry of Science and Technology of the People's Republic of China; Fundamental Research Funds for the Central Universities; Foundation for Innovative Research Groups of the National Natural Science Foundation of China; State Key Laboratory of Software Development Environment","keywords":"Computer science; Correctness; Question answering; Bottleneck; Knowledge base; Relation (database); Information bottleneck method; Artificial intelligence; Information retrieval; Machine learning; Data mining; Programming language","score_opus":0.018027592806274667,"score_gpt":0.2547266321229922,"score_spread":0.2366990393167175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973947483","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012261815,0.0009982993,0.9774239,0.001435195,0.00006566902,0.00027491772,0.0007667118,0.0039333478,0.0028401662],"genre_scores_gemma":[0.26572996,0.001421546,0.7230825,0.0010210393,0.00019142579,0.000554246,0.0038261947,0.00034496657,0.003828107],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99337476,0.0033689274,0.00051321,0.0012983357,0.0011153579,0.00032934893],"domain_scores_gemma":[0.98118275,0.01415622,0.0005231279,0.0024763388,0.0013905993,0.00027101656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064654374,0.0009435865,0.0012868525,0.0026028154,0.0011541437,0.004231533,0.0033049895,0.002440982,0.005597431],"category_scores_gemma":[0.032495465,0.0010740019,0.0018905847,0.002651701,0.0018891797,0.010184543,0.004265396,0.0030027267,0.0021085145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046597264,0.00049568055,0.0037484702,0.0019502934,0.00036464696,0.00040445197,0.0023689603,0.16484994,0.00937726,0.22465442,0.023069376,0.56825054],"study_design_scores_gemma":[0.00004903931,0.00008291614,0.00065288116,0.00017133157,0.000100631354,0.00019718752,0.0003307926,0.6059222,0.0067227096,0.36130217,0.024428392,0.000039706774],"about_ca_topic_score_codex":0.007240842,"about_ca_topic_score_gemma":0.006052153,"teacher_disagreement_score":0.007240842,"about_ca_system_score_codex":0.0016209054,"about_ca_system_score_gemma":0.0018694912,"threshold_uncertainty_score":0.03419292},"labels":[],"label_agreement":null},{"id":"W2974383016","doi":"10.48550/arxiv.1909.07512","title":"Short-Text Classification Using Unsupervised Keyword Expansion","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Sentence; Word (group theory); Artificial intelligence; Natural language processing; Feature (linguistics); Limiting; Process (computing); Speech recognition; Linguistics","score_opus":0.196499995122714,"score_gpt":0.21952941962633857,"score_spread":0.02302942450362458,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2974383016","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2174885,0.0010249713,0.7652746,0.00031570494,0.00024281458,0.00074616156,0.001874903,0.009675572,0.0033568114],"genre_scores_gemma":[0.5811419,0.00044562126,0.40075904,0.00021066006,0.00031711036,0.0008886653,0.008168033,0.00048234823,0.007586699],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991013,0.00026434707,0.00008876177,0.00033500773,0.00015192252,0.00005864373],"domain_scores_gemma":[0.99623543,0.002020989,0.0003767376,0.00045608386,0.00080475066,0.00010608429],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011314776,0.0009842522,0.00066034327,0.0018254799,0.00034096366,0.0006168058,0.0010084434,0.0007696468,0.0018121318],"category_scores_gemma":[0.0040843887,0.00029435763,0.0009050278,0.0014122853,0.00042446543,0.0020813101,0.0007698522,0.000996022,0.0020062479],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010208907,0.0004761655,0.0069298553,0.0004587723,0.00016349171,0.00043754026,0.00064779504,0.0841071,0.08408882,0.0037612617,0.011593122,0.8063152],"study_design_scores_gemma":[0.00005442341,0.00032526476,0.0026693542,0.000024477085,0.00005049295,0.00028025237,0.00013809818,0.9588663,0.027174992,0.0056099882,0.0047659916,0.000040439147],"about_ca_topic_score_codex":0.0014790891,"about_ca_topic_score_gemma":0.0022757624,"teacher_disagreement_score":0.0018254799,"about_ca_system_score_codex":0.0005130588,"about_ca_system_score_gemma":0.00061397697,"threshold_uncertainty_score":0.00606215},"labels":[],"label_agreement":null},{"id":"W2974904219","doi":"10.48550/arxiv.1909.08089","title":"Extractive Summarization of Long Documents by Combining Global and Local Context","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Context (archaeology); Computer science; Information retrieval; ROUGE; Data science; Artificial intelligence; Natural language processing; History; Archaeology","score_opus":0.039371286163352094,"score_gpt":0.19511875905608572,"score_spread":0.15574747289273363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2974904219","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036165282,0.008626125,0.94699085,0.0008101133,0.000298911,0.00017458135,0.0012030201,0.0035644071,0.002166657],"genre_scores_gemma":[0.46715114,0.005700965,0.5025944,0.0006000956,0.0019155066,0.00058035325,0.0073124357,0.0006808797,0.013464226],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991485,0.00018318706,0.00008821061,0.0003251225,0.00020200641,0.00005294133],"domain_scores_gemma":[0.9979729,0.00090560503,0.0003265664,0.0002451203,0.00048085957,0.00006898412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013087216,0.0018649751,0.0014086765,0.0029990033,0.0005147051,0.0017571428,0.0015436303,0.0011811176,0.001680281],"category_scores_gemma":[0.004358792,0.0005242547,0.0010705272,0.0026426492,0.0004608488,0.0034350387,0.0010755084,0.0015037518,0.0016956773],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004809397,0.00020485642,0.0033912887,0.00089018507,0.0003835359,0.00030428622,0.0005357555,0.084304646,0.028358871,0.0052929944,0.011501124,0.8643516],"study_design_scores_gemma":[0.00006397689,0.00042835696,0.0035233886,0.00012374892,0.00051604124,0.0002522981,0.00019840382,0.9377657,0.0193971,0.01921819,0.018432217,0.00008054908],"about_ca_topic_score_codex":0.0028867002,"about_ca_topic_score_gemma":0.0074056596,"teacher_disagreement_score":0.0029990033,"about_ca_system_score_codex":0.000711608,"about_ca_system_score_gemma":0.0010121978,"threshold_uncertainty_score":0.0069212914},"labels":[],"label_agreement":null},{"id":"W2976855657","doi":"10.1007/s10791-019-09364-x","title":"ReBoost: a retrieval-boosted sequence-to-sequence model for neural response generation","year":2019,"lang":"en","type":"article","venue":"Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Computer science; Benchmark (surveying); Conversation; Sequence (biology); Process (computing); Artificial intelligence; Natural language generation; Artificial neural network; Language model; Natural language processing; Recurrent neural network; Machine learning; Natural language; Programming language; Linguistics","score_opus":0.06926171490965688,"score_gpt":0.29959857846808324,"score_spread":0.23033686355842636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2976855657","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0124847805,0.0016706721,0.9764766,0.00040270877,0.00027719451,0.00018415894,0.00047849747,0.007150414,0.000874906],"genre_scores_gemma":[0.3098129,0.0012530565,0.6651555,0.0011058195,0.0005669449,0.00093753287,0.0032643736,0.0012766805,0.016627235],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989858,0.00036090385,0.00005319845,0.00025191015,0.00020676399,0.00014137538],"domain_scores_gemma":[0.9986563,0.00070633215,0.00007757792,0.000123855,0.00035486696,0.00008095574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023317148,0.0015646098,0.002680978,0.001497456,0.00065042154,0.0009734162,0.0042623007,0.0029234914,0.004797721],"category_scores_gemma":[0.0041233106,0.0007834593,0.0014742132,0.0017304125,0.00058068056,0.0016675759,0.0013692753,0.0033647686,0.0029566728],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009091464,0.0006341359,0.0007158202,0.00036339313,0.00033575587,0.00011762115,0.00008754877,0.28351018,0.012132262,0.005943159,0.02459648,0.67065454],"study_design_scores_gemma":[0.000021823107,0.000052857347,0.00008452019,0.000008001718,0.000017480377,0.000018715096,0.000004084248,0.9949279,0.0017742163,0.00224053,0.0008398411,0.000010054155],"about_ca_topic_score_codex":0.011684983,"about_ca_topic_score_gemma":0.013936229,"teacher_disagreement_score":0.011684983,"about_ca_system_score_codex":0.0010987844,"about_ca_system_score_gemma":0.0016416564,"threshold_uncertainty_score":0.02323395},"labels":[],"label_agreement":null},{"id":"W2978238777","doi":"10.22215/etd/2019-13523","title":"Data-Driven Creativity Enhancement Through Word Association","year":2019,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Creativity; Paragraph; Word Association; Word (group theory); Association (psychology); Task (project management); Blank; Computer science; Selection (genetic algorithm); Scale (ratio); Natural language processing; Artificial intelligence; Psychology; Linguistics; Engineering; Social psychology; World Wide Web; Geography","score_opus":0.05383212939535771,"score_gpt":0.32749862881887365,"score_spread":0.27366649942351595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2978238777","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12876002,0.0008708034,0.86147577,0.0005767579,0.00013871347,0.00027847837,0.0010988141,0.0042188605,0.0025818623],"genre_scores_gemma":[0.46111265,0.0005000491,0.5304557,0.00012945205,0.00017685074,0.00043433,0.003713201,0.00040313104,0.0030746653],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982431,0.00057321286,0.00017568933,0.0005631105,0.00035878544,0.000086221255],"domain_scores_gemma":[0.9907969,0.0063853986,0.0005818306,0.00079107,0.0011839435,0.00026076793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024184023,0.00091956864,0.00096283754,0.0029355155,0.00072961004,0.002349421,0.0013701447,0.0008133197,0.0017549469],"category_scores_gemma":[0.013000529,0.0006147399,0.0011569302,0.0036242364,0.0005570812,0.0029694438,0.0018831043,0.0015294731,0.0016067527],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001018864,0.0007133818,0.019752333,0.0005967442,0.00028532688,0.0002545671,0.0017923054,0.06738534,0.031556867,0.0051202048,0.008947403,0.86257666],"study_design_scores_gemma":[0.000087489454,0.0001685204,0.003914605,0.00005685755,0.00014833658,0.0001629233,0.0005127553,0.95064795,0.02082882,0.01364854,0.009766839,0.000056410885],"about_ca_topic_score_codex":0.0016868872,"about_ca_topic_score_gemma":0.0032048924,"teacher_disagreement_score":0.0029355155,"about_ca_system_score_codex":0.00048347356,"about_ca_system_score_gemma":0.0009057882,"threshold_uncertainty_score":0.0127898455},"labels":[],"label_agreement":null},{"id":"W2980073313","doi":"10.1162/coli_a_00363","title":"Scalable Micro-planned Generation of Discourse from Structured Data","year":2019,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Natural language processing; Interpretability; Scalability; Artificial intelligence; Natural language generation; Sentence; Pipeline (software); Paragraph; Fluency; Natural language understanding; Data manipulation language; Natural language; Information retrieval; Programming language; Database; World Wide Web","score_opus":0.06693568368884965,"score_gpt":0.266282299033941,"score_spread":0.19934661534509135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2980073313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010048107,0.00021847773,0.9574152,0.00036356243,0.0001088448,0.0002812419,0.0037777475,0.025406878,0.0023798968],"genre_scores_gemma":[0.0986937,0.00020697914,0.88215744,0.00013832642,0.00007518547,0.0004002144,0.013432839,0.0018557154,0.0030395284],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987833,0.00038092418,0.00009299593,0.0004314519,0.0002597541,0.000051534687],"domain_scores_gemma":[0.99642485,0.0021671408,0.00019716249,0.00057319965,0.0005316508,0.000105923486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016967893,0.0009946523,0.00071139814,0.0015977937,0.00064963446,0.0015960656,0.0015870837,0.00076307374,0.0075505455],"category_scores_gemma":[0.008035829,0.0005589837,0.0012363602,0.0012222787,0.000643528,0.002357733,0.0019695088,0.001125186,0.004113657],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006137777,0.0002780102,0.002990936,0.0015448525,0.0001537692,0.0010144423,0.0028633464,0.046189774,0.054995626,0.054939356,0.0722007,0.7622154],"study_design_scores_gemma":[0.00012854453,0.00013320197,0.0008206692,0.00010367055,0.00007517619,0.00036264205,0.0006722469,0.79806155,0.06775742,0.054715887,0.0770875,0.000081507394],"about_ca_topic_score_codex":0.0026950417,"about_ca_topic_score_gemma":0.0043609035,"teacher_disagreement_score":0.0075505455,"about_ca_system_score_codex":0.0007689247,"about_ca_system_score_gemma":0.0018777612,"threshold_uncertainty_score":0.025259078},"labels":[],"label_agreement":null},{"id":"W2980386803","doi":"10.1007/978-3-030-33509-0_13","title":"Introducing Connotation Similarity","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Connotation; Similarity (geometry); Cosine similarity; Semantic similarity; Natural language processing; Meaning (existential); Context (archaeology); Computer science; Artificial intelligence; Information retrieval; Linguistics; Pattern recognition (psychology); Psychology; History","score_opus":0.017260202823701925,"score_gpt":0.22344368508283416,"score_spread":0.20618348225913224,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2980386803","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0076345117,0.0057871854,0.79234606,0.0046524187,0.0043088715,0.00018815405,0.00086738274,0.0030886498,0.18112685],"genre_scores_gemma":[0.23717833,0.0049624867,0.5651263,0.0025232774,0.0045305826,0.0005982104,0.0029595701,0.0050792987,0.17704196],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99592316,0.0013347853,0.0003528965,0.0011844819,0.000990906,0.0002138589],"domain_scores_gemma":[0.9958598,0.0017315351,0.00017584077,0.0010333346,0.001009877,0.0001895759],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021399325,0.001063786,0.0012390197,0.0064656734,0.0029355797,0.0069762478,0.001611523,0.0016992084,0.045901638],"category_scores_gemma":[0.015332068,0.0007710338,0.0014074323,0.007981593,0.0031791348,0.016502013,0.0061071496,0.0031526634,0.014256748],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003303333,0.000020339972,0.00033009352,0.00013357987,0.000019178926,0.000089105655,0.001035626,0.0004458928,0.0008237654,0.8019913,0.025920048,0.1691581],"study_design_scores_gemma":[0.000006246342,0.000017518048,0.00027903714,0.00012880472,0.000021870806,0.00036946515,0.0004779518,0.0059694042,0.0010513486,0.6971075,0.29453626,0.000034601457],"about_ca_topic_score_codex":0.0018468624,"about_ca_topic_score_gemma":0.0019528703,"teacher_disagreement_score":0.045901638,"about_ca_system_score_codex":0.0017027128,"about_ca_system_score_gemma":0.0010730185,"threshold_uncertainty_score":0.15355623},"labels":[],"label_agreement":null},{"id":"W2980839612","doi":"10.2196/14850","title":"Combining Contextualized Embeddings and Prior Knowledge for Clinical Named Entity Recognition: Evaluation Study","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Eli Lilly and Company","keywords":"Computer science; Named-entity recognition; Natural language processing; Artificial intelligence; Word embedding; Lexicon; F1 score; Deep learning; Context (archaeology); Leverage (statistics); Embedding; Information retrieval; Task (project management)","score_opus":0.11758840970701676,"score_gpt":0.4348812763784703,"score_spread":0.3172928666714535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2980839612","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92393744,0.022087976,0.031323086,0.001178936,0.001077869,0.0012746991,0.0070508113,0.0040190904,0.00805011],"genre_scores_gemma":[0.92909527,0.0045949304,0.038670927,0.0004609654,0.00036594298,0.0004750418,0.023895752,0.00014647807,0.002294736],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959596,0.0016323898,0.0005587628,0.0008575482,0.00078881526,0.00020291003],"domain_scores_gemma":[0.98965746,0.0060902443,0.0005528278,0.0011683047,0.0019907902,0.0005404201],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072852382,0.0022790323,0.0015396212,0.0019862745,0.0005132755,0.0011460105,0.0014234796,0.002083156,0.0025292882],"category_scores_gemma":[0.01684155,0.00035706337,0.0013216843,0.001269957,0.000652903,0.0028517002,0.0017375487,0.001063901,0.001446188],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008589764,0.004758471,0.06948738,0.0038789408,0.0027960099,0.000932178,0.00034660468,0.059737574,0.009697278,0.00072910456,0.034374237,0.80467254],"study_design_scores_gemma":[0.0013477259,0.010216838,0.061471608,0.00054390245,0.002959608,0.0024157139,0.00063694245,0.8827564,0.022323646,0.0021061038,0.012958074,0.0002634162],"about_ca_topic_score_codex":0.0060814233,"about_ca_topic_score_gemma":0.0071409526,"teacher_disagreement_score":0.0072852382,"about_ca_system_score_codex":0.0010813855,"about_ca_system_score_gemma":0.0011924505,"threshold_uncertainty_score":0.038528502},"labels":[],"label_agreement":null},{"id":"W2981198292","doi":"10.48550/arxiv.1910.06575","title":"Aligning Cross-Lingual Entities with Multi-Aspect Information","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of Waterloo","funders":"","keywords":"Computer science; Exploit; Entity linking; ENCODE; Embedding; Natural language processing; Artificial intelligence; Task (project management); Benchmark (surveying); Vector space; Graph; Bridge (graph theory); Information retrieval; Theoretical computer science; Knowledge base","score_opus":0.06723116264791881,"score_gpt":0.1982985583223212,"score_spread":0.1310673956744024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2981198292","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15314049,0.0019103885,0.8131547,0.000911288,0.00028587368,0.0001741589,0.005936199,0.013374414,0.0111124115],"genre_scores_gemma":[0.707069,0.0009825994,0.25610802,0.00040083742,0.00014816382,0.00011793221,0.02549329,0.0011153816,0.008564747],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991154,0.00014947362,0.00005125827,0.0004773219,0.000122633,0.00008385693],"domain_scores_gemma":[0.9988545,0.0003036352,0.00013757174,0.00046493098,0.00018245695,0.00005696752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00094229676,0.0011978939,0.00066721754,0.0033989311,0.0006848234,0.0015263625,0.0014722617,0.0010304683,0.0023188055],"category_scores_gemma":[0.0029384173,0.00048379213,0.0013525081,0.004982258,0.0004982183,0.005609752,0.002582514,0.001608978,0.0019012486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041742212,0.00034970912,0.026972352,0.00046821477,0.0005570234,0.00071242725,0.00087144796,0.07391146,0.023258485,0.027832542,0.032631215,0.8120176],"study_design_scores_gemma":[0.000026338932,0.00006779649,0.009754348,0.00007355984,0.00022609803,0.00046141728,0.00055117655,0.88018125,0.017830316,0.04916719,0.041605502,0.000054990996],"about_ca_topic_score_codex":0.0099698175,"about_ca_topic_score_gemma":0.02725535,"teacher_disagreement_score":0.0099698175,"about_ca_system_score_codex":0.0009879194,"about_ca_system_score_gemma":0.0009232704,"threshold_uncertainty_score":0.01982361},"labels":[],"label_agreement":null},{"id":"W2982424689","doi":"10.1016/j.yjbinx.2019.100057","title":"A survey of word embeddings for clinical text","year":2019,"lang":"en","type":"review","venue":"Journal of Biomedical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":240,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto; St. Michael's Hospital","funders":"","keywords":"Word (group theory); Computer science; Blueprint; Natural language processing; Artificial intelligence; Information retrieval; Data science; Linguistics","score_opus":0.26105588914672456,"score_gpt":0.4618205473794558,"score_spread":0.20076465823273126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2982424689","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032106903,0.93502265,0.050901737,0.0026788164,0.00096855283,0.00016551832,0.0013133879,0.0005181994,0.005220424],"genre_scores_gemma":[0.023398766,0.91858184,0.048522267,0.0009238912,0.0011020483,0.00035847758,0.0031669554,0.00021264923,0.003733077],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990318,0.00034349947,0.0001510125,0.00020248214,0.00023714025,0.000034090746],"domain_scores_gemma":[0.99611413,0.0028826206,0.00023256356,0.000178641,0.0005324529,0.00005961197],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021837063,0.0016232847,0.0012353355,0.0046125576,0.00024269287,0.0015714959,0.0011782959,0.0010017267,0.0043637515],"category_scores_gemma":[0.009502555,0.00048495835,0.0009553081,0.0046832273,0.0006978462,0.004003488,0.0011656638,0.0013442213,0.0034613097],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003651737,0.00004447413,0.000471043,0.0059984396,0.000064435815,0.000023340777,0.000080237645,0.00093816844,0.0007137039,0.0056957984,0.01524314,0.9706906],"study_design_scores_gemma":[0.00005193171,0.0004418559,0.0074137924,0.016037855,0.00040731262,0.0017677625,0.0005956394,0.01583408,0.006046632,0.04248301,0.90871096,0.00020913775],"about_ca_topic_score_codex":0.0013813792,"about_ca_topic_score_gemma":0.0013007518,"teacher_disagreement_score":0.0046125576,"about_ca_system_score_codex":0.0006663167,"about_ca_system_score_gemma":0.0014393945,"threshold_uncertainty_score":0.014598191},"labels":[],"label_agreement":null},{"id":"W2982709062","doi":"","title":"Ordered Memory","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Inference; Task (project management); Recurrent neural network; Representation (politics); Artificial intelligence; Tree (set theory); Theoretical computer science; Machine learning; Artificial neural network; Mathematics","score_opus":0.0850267741566962,"score_gpt":0.17822951988973412,"score_spread":0.09320274573303793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2982709062","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.078293994,0.0019050959,0.8849045,0.0009861729,0.00050375273,0.00009845664,0.0019261623,0.0072834585,0.024098452],"genre_scores_gemma":[0.82420015,0.0012253753,0.13034451,0.00059082924,0.00019987617,0.00017880394,0.002985758,0.00050092646,0.03977379],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9997962,0.000022993267,0.000014368689,0.0000859719,0.000038203794,0.000042114563],"domain_scores_gemma":[0.999584,0.0001035232,0.00004684258,0.0001383122,0.00009795174,0.00002935222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00028842242,0.00087741105,0.0005737065,0.0004832918,0.00037919547,0.0011121649,0.001949701,0.00078157603,0.012257856],"category_scores_gemma":[0.0014708443,0.0003496809,0.00069915026,0.00064566184,0.00054435874,0.0033500094,0.0010045235,0.0010857023,0.0034189636],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005629759,0.00023695594,0.00392145,0.00044156192,0.00017759055,0.00047447713,0.0002937253,0.18668334,0.02422174,0.14580348,0.030757023,0.60642564],"study_design_scores_gemma":[0.000025406782,0.000112176334,0.00055932836,0.000034676184,0.00005947541,0.00015865524,0.00004637227,0.8752724,0.009634213,0.10208478,0.011983102,0.000029382994],"about_ca_topic_score_codex":0.0042100507,"about_ca_topic_score_gemma":0.0090949675,"teacher_disagreement_score":0.012257856,"about_ca_system_score_codex":0.0005866197,"about_ca_system_score_gemma":0.00075813883,"threshold_uncertainty_score":0.041006565},"labels":[],"label_agreement":null},{"id":"W2983102021","doi":"10.48550/arxiv.1911.06136","title":"KEPLER: A Unified Model for Knowledge Embedding and Pre-trained Language Representation","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; Université de Montréal","funders":"","keywords":"Kepler; Embedding; Computer science; Benchmark (surveying); Representation (politics); Language model; Natural language processing; Construct (python library); ENCODE; Artificial intelligence; Programming language","score_opus":0.10695254559882145,"score_gpt":0.25321215459734564,"score_spread":0.1462596089985242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2983102021","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009384645,0.0007936897,0.9788791,0.0005578476,0.00011799631,0.00013690736,0.0021927373,0.0062351897,0.0017018284],"genre_scores_gemma":[0.29223028,0.0016202643,0.6680396,0.000943237,0.00023559334,0.0010182788,0.023676056,0.0011039384,0.011132755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99889165,0.00033014241,0.0000900896,0.00043992655,0.00015393601,0.00009423786],"domain_scores_gemma":[0.9974584,0.0012989058,0.00017670989,0.0006199466,0.00034882812,0.0000972284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015750148,0.0017174337,0.001004169,0.0021573408,0.0005274568,0.0018892231,0.0033419984,0.0019216266,0.0036176543],"category_scores_gemma":[0.008708905,0.00083370326,0.0018310761,0.0023024138,0.0007918612,0.006954987,0.002941407,0.0038881663,0.0034370064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030689544,0.00031210316,0.0025175102,0.00054112,0.00030322993,0.00031974833,0.00043430857,0.3517525,0.005892227,0.034059115,0.032781105,0.57078016],"study_design_scores_gemma":[0.000017087874,0.000039050647,0.00020996277,0.000035602236,0.0000372366,0.00007311592,0.000038796967,0.9699443,0.0019324443,0.022601,0.0050499546,0.000021375226],"about_ca_topic_score_codex":0.006241342,"about_ca_topic_score_gemma":0.009962868,"teacher_disagreement_score":0.006241342,"about_ca_system_score_codex":0.0012271397,"about_ca_system_score_gemma":0.0016592424,"threshold_uncertainty_score":0.012410045},"labels":[],"label_agreement":null},{"id":"W2983647115","doi":"10.1109/ijcnn52387.2021.9533769","title":"Soft-Label Dataset Distillation and Text Dataset Distillation","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Distillation; Computer science; Code (set theory); Artificial intelligence; Sample (material); Pattern recognition (psychology); Image (mathematics); Task (project management); MNIST database; Machine learning; Data mining; Deep learning; Chromatography; Engineering","score_opus":0.048223260901226375,"score_gpt":0.29652199813914903,"score_spread":0.24829873723792267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2983647115","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05894656,0.0011967672,0.8594255,0.0016513623,0.00078730524,0.0008977334,0.011237491,0.060127813,0.005729442],"genre_scores_gemma":[0.1925547,0.00029710334,0.76209635,0.0012033253,0.00024732194,0.0013680765,0.03219669,0.003056865,0.0069795027],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971535,0.00073929277,0.0002162244,0.0010820874,0.00057223183,0.00023658246],"domain_scores_gemma":[0.9937232,0.0020990632,0.0003331415,0.0029323457,0.00069872424,0.00021362209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029581152,0.0025973828,0.001565923,0.002415835,0.0012575674,0.0024002164,0.0033239613,0.002108629,0.012232077],"category_scores_gemma":[0.015612869,0.00091802917,0.0022940286,0.0028157816,0.0016552213,0.006112138,0.0047876956,0.0041278163,0.0056430935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013376873,0.0006740743,0.0051581734,0.001049313,0.00029590912,0.00038358365,0.0004761372,0.11807498,0.042482525,0.023743324,0.07721219,0.7291121],"study_design_scores_gemma":[0.00024542687,0.00025249415,0.0014613846,0.00007739318,0.0000640729,0.00029684394,0.00016796269,0.88390523,0.043140605,0.03585914,0.03443351,0.00009592337],"about_ca_topic_score_codex":0.003719966,"about_ca_topic_score_gemma":0.008734725,"teacher_disagreement_score":0.012232077,"about_ca_system_score_codex":0.0016835682,"about_ca_system_score_gemma":0.0020660462,"threshold_uncertainty_score":0.040920317},"labels":[],"label_agreement":null},{"id":"W2984469754","doi":"10.18653/v1/d19-1540","title":"Bridging the Gap between Relevance Matching and Semantic Matching for Short Text Similarity Modeling","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":71,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Bridging (networking); Computer science; Natural language processing; Relevance (law); Matching (statistics); Similarity (geometry); Semantic similarity; Artificial intelligence; Semantic matching; Natural language; Information retrieval; Mathematics; Image (mathematics)","score_opus":0.04279903629725576,"score_gpt":0.2781462756295075,"score_spread":0.23534723933225174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2984469754","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03484139,0.006938775,0.95240176,0.001087304,0.0002824907,0.0002150708,0.00042732165,0.0010408495,0.0027650993],"genre_scores_gemma":[0.6070431,0.0037688597,0.38166362,0.00059388607,0.00096540654,0.00045552631,0.0023255881,0.00037953505,0.0028045394],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9907301,0.004534072,0.00086047093,0.0016371952,0.0018468513,0.0003913621],"domain_scores_gemma":[0.9864342,0.008627483,0.00056815456,0.0025049956,0.0015119625,0.00035313165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007845947,0.0008192507,0.0020244871,0.004909308,0.0012895813,0.003263999,0.0021408673,0.002613018,0.0031525746],"category_scores_gemma":[0.02830907,0.0005592899,0.0016315449,0.0048101223,0.001429915,0.011839579,0.004581454,0.0019486581,0.0017471943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010961916,0.00060463825,0.004810032,0.0011481526,0.0004462983,0.00028544615,0.0012857724,0.02638025,0.011419517,0.10893182,0.017064461,0.82652736],"study_design_scores_gemma":[0.000103726365,0.00033149763,0.0027462419,0.00016345644,0.0002296412,0.0003840628,0.0005204577,0.71634066,0.0067823343,0.25437704,0.017928824,0.000092034694],"about_ca_topic_score_codex":0.003089031,"about_ca_topic_score_gemma":0.0034123783,"teacher_disagreement_score":0.007845947,"about_ca_system_score_codex":0.00096521823,"about_ca_system_score_gemma":0.0018827871,"threshold_uncertainty_score":0.041493833},"labels":[],"label_agreement":null},{"id":"W2984811147","doi":"10.18653/v1/d19-5627","title":"Transformer and seq2seq model for Paraphrase Generation","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Alberta Innovates; University of Lethbridge","keywords":"Paraphrase; Computer science; Transformer; Encoder; Sentence; Artificial intelligence; Natural language processing; Decoding methods; Speech recognition; Algorithm; Engineering; Voltage; Electrical engineering","score_opus":0.045415728540805424,"score_gpt":0.2551178433031982,"score_spread":0.20970211476239278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2984811147","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03088171,0.0014499204,0.9443419,0.00080876,0.00029886418,0.00035337027,0.0032183784,0.012631321,0.0060157366],"genre_scores_gemma":[0.5118039,0.0011085198,0.45395947,0.00081946264,0.00020344477,0.0007137937,0.011312112,0.0010256998,0.019053577],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995277,0.00013393817,0.000040833802,0.00015648465,0.00009498273,0.00004610839],"domain_scores_gemma":[0.9993771,0.00022955294,0.00004120488,0.0001223688,0.00018636465,0.000043475065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009922349,0.0012897047,0.00076135417,0.0010110668,0.00035937843,0.0008528596,0.002016245,0.0013259333,0.007682535],"category_scores_gemma":[0.0024307207,0.00043039274,0.0012785164,0.0009209242,0.0003989639,0.0022605439,0.0009642885,0.0017107682,0.0040491675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005973261,0.00041072816,0.0028060048,0.0006236779,0.00026133997,0.0008930291,0.00035709018,0.2655666,0.035474703,0.024927935,0.03998202,0.62809956],"study_design_scores_gemma":[0.000025980913,0.000089032634,0.00026892373,0.000011532975,0.000030730465,0.00018075341,0.000028906565,0.9760957,0.008990348,0.010202747,0.004058264,0.000017197606],"about_ca_topic_score_codex":0.005477617,"about_ca_topic_score_gemma":0.011017508,"teacher_disagreement_score":0.007682535,"about_ca_system_score_codex":0.0008129203,"about_ca_system_score_gemma":0.0011701109,"threshold_uncertainty_score":0.025700629},"labels":[],"label_agreement":null},{"id":"W2985067290","doi":"10.18653/v1/d19-1459","title":"Taskmaster-1: Toward a Realistic and Diverse Dialog Dataset","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":139,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Dialog box; Joint (building); Computer science; Natural language; Natural (archaeology); Natural language processing; Artificial intelligence; Engineering; World Wide Web; History; Architectural engineering; Archaeology","score_opus":0.0398766274616048,"score_gpt":0.25229376539206205,"score_spread":0.21241713793045724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2985067290","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1150854,0.0036316463,0.033015534,0.005336209,0.0021275505,0.002823171,0.80075604,0.019028543,0.018195977],"genre_scores_gemma":[0.072300576,0.00029535644,0.03181538,0.000909049,0.0002347444,0.0017739075,0.88491774,0.00044478694,0.0073085013],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962,0.001811991,0.00028380932,0.0008223338,0.0005762173,0.0003055709],"domain_scores_gemma":[0.99394506,0.0017071288,0.00030455575,0.0018167942,0.0011793514,0.0010472616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004290163,0.0018005908,0.0010875164,0.0022904149,0.0022812288,0.0019423807,0.0035598034,0.0036581308,0.008901863],"category_scores_gemma":[0.011635308,0.0006755318,0.0015681678,0.0018190787,0.00076244754,0.0033292372,0.00415392,0.0029681595,0.014411755],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016103315,0.0019932406,0.010761015,0.0016994505,0.00033153672,0.00055878336,0.0011072375,0.008077226,0.0051833573,0.0035456172,0.90863085,0.056501437],"study_design_scores_gemma":[0.0017333437,0.0015255839,0.04266193,0.0006512259,0.00029806822,0.001817466,0.004738601,0.108111426,0.012815612,0.013216303,0.8119903,0.0004401417],"about_ca_topic_score_codex":0.016204262,"about_ca_topic_score_gemma":0.031631663,"teacher_disagreement_score":0.016204262,"about_ca_system_score_codex":0.0015579602,"about_ca_system_score_gemma":0.0022017362,"threshold_uncertainty_score":0.032219887},"labels":[],"label_agreement":null},{"id":"W2985573020","doi":"10.18653/v1/d19-5810","title":"Cross-Task Knowledge Transfer for Query-Based Text Summarization","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Leverage (statistics); Natural language processing; Machine translation; Artificial intelligence; Task (project management); Multi-document summarization; Transfer of learning; Sentence; Information retrieval","score_opus":0.020091061961844847,"score_gpt":0.27673004615301783,"score_spread":0.256638984191173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2985573020","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30041242,0.0051967264,0.61633146,0.0033626554,0.000887251,0.0008437585,0.0051665856,0.044076513,0.023722645],"genre_scores_gemma":[0.84887165,0.00059644884,0.1298277,0.00061953656,0.0004075403,0.0004056906,0.009921646,0.0006666202,0.00868333],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970004,0.0013406391,0.00019059365,0.0009613351,0.00026100202,0.00024591415],"domain_scores_gemma":[0.98965317,0.0064509083,0.00040720438,0.002258795,0.0009775558,0.00025238626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047810315,0.0017576857,0.0011170257,0.0013591251,0.0007748769,0.0017913159,0.0017346704,0.0014351712,0.0063821524],"category_scores_gemma":[0.018997828,0.0003160284,0.0010531372,0.0012377428,0.000636301,0.004786257,0.002736376,0.0025007948,0.003939035],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009733339,0.0016722591,0.0054146587,0.00107414,0.00053780706,0.00049536553,0.0011617375,0.084173344,0.043777637,0.005066833,0.0362935,0.8193595],"study_design_scores_gemma":[0.00014559872,0.00093495805,0.0036996438,0.000051323066,0.0002457528,0.00020145676,0.00042603843,0.91296816,0.05590197,0.013647318,0.011688292,0.000089509784],"about_ca_topic_score_codex":0.0051601157,"about_ca_topic_score_gemma":0.006831586,"teacher_disagreement_score":0.0063821524,"about_ca_system_score_codex":0.0011972671,"about_ca_system_score_gemma":0.0015044843,"threshold_uncertainty_score":0.025284827},"labels":[],"label_agreement":null},{"id":"W2985678088","doi":"10.18653/v1/d19-5513","title":"Contextualized Word Representations from Distant Supervision with and for NER","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Nvidia","keywords":"Computer science; Representation (politics); Word (group theory); Natural language processing; Named-entity recognition; Artificial intelligence; Linguistics; Task (project management)","score_opus":0.032033742000887405,"score_gpt":0.2866846165628594,"score_spread":0.254650874561972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2985678088","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035125718,0.002192051,0.9321127,0.00048036507,0.00037712595,0.00016114816,0.0046978304,0.017193483,0.007659565],"genre_scores_gemma":[0.5144009,0.0013831398,0.4437105,0.00055389863,0.0003613385,0.00039900537,0.025501665,0.0014316676,0.0122578945],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999295,0.00019087241,0.000042817777,0.00030107994,0.000096338175,0.000073884905],"domain_scores_gemma":[0.99912304,0.00017126925,0.000068181515,0.00048366957,0.00011500031,0.000038879454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006811059,0.0015224019,0.00095481315,0.0010797761,0.0004556457,0.000801661,0.0011901684,0.00109233,0.0069812117],"category_scores_gemma":[0.0028879507,0.00030644314,0.0007988263,0.0012955737,0.0005065316,0.0038897975,0.0020273959,0.0018602437,0.0045149163],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042403617,0.0002699419,0.0020612965,0.0006465157,0.00018134955,0.00022287875,0.0002767802,0.05819198,0.02651782,0.023650173,0.043161843,0.8443953],"study_design_scores_gemma":[0.00008408206,0.00037169078,0.0024696707,0.00021011285,0.00021310887,0.0004897413,0.0002798154,0.75926054,0.036686387,0.1298868,0.06993697,0.0001111452],"about_ca_topic_score_codex":0.0023996946,"about_ca_topic_score_gemma":0.0070776083,"teacher_disagreement_score":0.0069812117,"about_ca_system_score_codex":0.00042434825,"about_ca_system_score_gemma":0.00091538345,"threshold_uncertainty_score":0.02335453},"labels":[],"label_agreement":null},{"id":"W2986143601","doi":"10.18653/v1/k19-1020","title":"Fully Unsupervised Crosslingual Semantic Textual Similarity Metric Based on BERT for Identifying Parallel Data","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Semantic similarity; Artificial intelligence; Metric (unit); Similarity (geometry); Natural language processing; Information retrieval; Pattern recognition (psychology)","score_opus":0.10503656253648687,"score_gpt":0.3355009182367941,"score_spread":0.23046435570030727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2986143601","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16984878,0.0012922236,0.8122032,0.00025802187,0.00015646937,0.00029843926,0.003983625,0.0058721025,0.0060870894],"genre_scores_gemma":[0.6830339,0.00031798164,0.29456794,0.00014681777,0.00010273234,0.00048645967,0.017095067,0.00063503935,0.0036139947],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974075,0.0007883518,0.00026543756,0.00073055364,0.000644188,0.00016402474],"domain_scores_gemma":[0.99643815,0.001012624,0.00040638872,0.00085848325,0.0011333287,0.00015104795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017982834,0.000985696,0.00081635645,0.0039168135,0.00071702217,0.0012283928,0.0011963061,0.0009406242,0.0023428292],"category_scores_gemma":[0.008854412,0.00023576934,0.0007517515,0.0028315547,0.00085441116,0.0034156456,0.0026028892,0.0011653202,0.0017192445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010321668,0.0005073724,0.018946959,0.0008793144,0.00030544781,0.0005035138,0.0010975309,0.057762373,0.0543652,0.02542922,0.018336618,0.8208343],"study_design_scores_gemma":[0.000060736933,0.00042506662,0.014065546,0.0000830379,0.00012982644,0.0008128729,0.0008203601,0.8811128,0.04305773,0.03774379,0.021568723,0.00011952612],"about_ca_topic_score_codex":0.0029895653,"about_ca_topic_score_gemma":0.005270708,"teacher_disagreement_score":0.0039168135,"about_ca_system_score_codex":0.0008179013,"about_ca_system_score_gemma":0.0011880983,"threshold_uncertainty_score":0.009510338},"labels":[],"label_agreement":null},{"id":"W2987386900","doi":"10.1016/j.asoc.2019.105913","title":"Structural block driven enhanced convolutional neural representation for relation extraction","year":2019,"lang":"en","type":"article","venue":"Applied Soft Computing","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"H2020 Marie Skłodowska-Curie Actions; Horizon 2020; Horizon 2020 Framework Programme","keywords":"Block (permutation group theory); Computer science; Convolutional neural network; Representation (politics); Extraction (chemistry); Relationship extraction; Pattern recognition (psychology); Relation (database); Artificial intelligence; Algorithm; Mathematics; Data mining; Chemistry; Chromatography; Combinatorics","score_opus":0.02200493358396193,"score_gpt":0.27983557893516414,"score_spread":0.2578306453512022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2987386900","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.072509766,0.0022118648,0.9096308,0.0004779271,0.00023242524,0.00014189332,0.003154618,0.00574491,0.005895829],"genre_scores_gemma":[0.6461361,0.001732597,0.3153279,0.00025982398,0.00026160348,0.00027275935,0.011756518,0.00035674818,0.02389589],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99972636,0.000040135314,0.000020544949,0.000082095205,0.00007531763,0.000055580505],"domain_scores_gemma":[0.9996055,0.00013295336,0.000037441892,0.000082046376,0.00012171003,0.000020480613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00035897217,0.0006634598,0.000625906,0.001559322,0.00029333172,0.00064860255,0.0009225953,0.0007585276,0.004526693],"category_scores_gemma":[0.0009983074,0.00023826961,0.0007794711,0.0020445648,0.00019516621,0.0010636051,0.00066125015,0.0009439969,0.0030470267],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031861535,0.00019822428,0.0017116797,0.00019365588,0.000095718795,0.00016174707,0.00009812673,0.039698422,0.04898277,0.009007179,0.014702414,0.8848315],"study_design_scores_gemma":[0.000010497774,0.00006167497,0.0014605104,0.000023503126,0.00006935823,0.000096339565,0.000026665772,0.9721897,0.013818641,0.0057782107,0.006452257,0.000012654705],"about_ca_topic_score_codex":0.008031946,"about_ca_topic_score_gemma":0.015166789,"teacher_disagreement_score":0.008031946,"about_ca_system_score_codex":0.000512851,"about_ca_system_score_gemma":0.0012467419,"threshold_uncertainty_score":0.01597035},"labels":[],"label_agreement":null},{"id":"W2987844692","doi":"10.18653/v1/d19-5607","title":"Transformer-based Model for Single Documents Neural Summarization","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Alberta Innovates; University of Lethbridge","keywords":"Automatic summarization; Computer science; Encoder; Transformer; ENCODE; Decoding methods; Artificial intelligence; Correctness; Natural language processing; Vocabulary; Speech recognition; Programming language; Algorithm","score_opus":0.031126324528247012,"score_gpt":0.25359788429103086,"score_spread":0.22247155976278385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2987844692","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015422851,0.0005144236,0.96897596,0.0002851522,0.00009739824,0.00009089179,0.0006761829,0.012220218,0.0017169251],"genre_scores_gemma":[0.42281526,0.0007559736,0.5553195,0.00042946739,0.00026391042,0.00033097333,0.004995552,0.0008447816,0.014244692],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99963725,0.000071643066,0.000030066862,0.00012705622,0.000090323716,0.0000436371],"domain_scores_gemma":[0.9993895,0.00021237975,0.00005653905,0.00012872112,0.0001780466,0.00003473544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008134173,0.00093011174,0.0008886398,0.0008347204,0.00033615547,0.00085691334,0.0016862346,0.00082922686,0.004393986],"category_scores_gemma":[0.0021269042,0.0003313456,0.0008407716,0.0009133762,0.00036179263,0.0023752453,0.00081392133,0.0014148544,0.0025955033],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005707752,0.00026227167,0.0008469302,0.0003307586,0.00016180046,0.00017753197,0.0002288222,0.1997805,0.05420168,0.013384401,0.01815513,0.71189946],"study_design_scores_gemma":[0.000024400491,0.0001046172,0.00014493291,0.000006727397,0.0000396186,0.000048267244,0.000023595721,0.9750244,0.015624471,0.005989623,0.002955597,0.000013715646],"about_ca_topic_score_codex":0.0055704615,"about_ca_topic_score_gemma":0.011527568,"teacher_disagreement_score":0.0055704615,"about_ca_system_score_codex":0.0009124724,"about_ca_system_score_gemma":0.0010817446,"threshold_uncertainty_score":0.0146993995},"labels":[],"label_agreement":null},{"id":"W2987869401","doi":"10.26615/978-954-452-056-4_018","title":"Learning Sentence Embeddings for Coherence Modelling and Beyond","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Coherence (philosophical gambling strategy); Sentence; Artificial intelligence; Heuristics; Natural language processing; Embedding","score_opus":0.0217580173074763,"score_gpt":0.2430688725556375,"score_spread":0.2213108552481612,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2987869401","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026039481,0.00044348263,0.9692527,0.00030757897,0.000078448604,0.000048720023,0.00055206125,0.0022553254,0.0010222221],"genre_scores_gemma":[0.53423166,0.00050816627,0.4571382,0.00026825356,0.00020292687,0.00022299125,0.0035422058,0.00064980896,0.0032357927],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993932,0.00025251182,0.00005663521,0.00017338326,0.00008393915,0.000040317867],"domain_scores_gemma":[0.9974746,0.0013584391,0.00032589666,0.00039126622,0.0003723053,0.00007743515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010034599,0.0009930317,0.00044970837,0.0010034197,0.00037172335,0.000948751,0.00082445185,0.000836932,0.0025772378],"category_scores_gemma":[0.0078008347,0.00032165635,0.00056154805,0.00092144596,0.00048353552,0.0036039746,0.001301776,0.0015945232,0.0011698841],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000532619,0.00031566713,0.0055709933,0.0007036893,0.00015391607,0.00038510005,0.0016999305,0.13952745,0.045652512,0.05340516,0.020333216,0.73171985],"study_design_scores_gemma":[0.00002431433,0.00013081537,0.0007652653,0.00004799966,0.00003117954,0.00009194273,0.00015204327,0.939236,0.008358434,0.043619126,0.0075179795,0.000024920879],"about_ca_topic_score_codex":0.0011980056,"about_ca_topic_score_gemma":0.0023480903,"teacher_disagreement_score":0.0025772378,"about_ca_system_score_codex":0.00031179254,"about_ca_system_score_gemma":0.00053761015,"threshold_uncertainty_score":0.008621752},"labels":[],"label_agreement":null},{"id":"W2990321249","doi":"10.1111/coin.12248","title":"A topic‐based term frequency normalization framework to enhance probabilistic information retrieval","year":2019,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wilfrid Laurier University; York University","funders":"Natural Sciences and Engineering Research Council of Canada; Yenepoya Research Centre","keywords":"Computer science; Divergence-from-randomness model; Normalization (sociology); Term Discrimination; Probabilistic logic; Artificial intelligence; Term (time); Language model; Natural language processing; Embedding; Sentence; Vector space model; Information retrieval; Visual Word","score_opus":0.017773968724452867,"score_gpt":0.2958569037642808,"score_spread":0.27808293503982795,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990321249","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009692645,0.002469305,0.984626,0.0002950042,0.00012342485,0.000075600656,0.00028975608,0.0014445924,0.0009836319],"genre_scores_gemma":[0.3959742,0.0033897408,0.5890187,0.00049652066,0.0008523762,0.0005979001,0.002178091,0.00047899596,0.007013365],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998185,0.000567945,0.00014621535,0.0004589811,0.00051087566,0.000130966],"domain_scores_gemma":[0.9981407,0.0007610609,0.00021062812,0.00027468597,0.0005555068,0.000057433677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031291246,0.0012012772,0.0015028807,0.0037507496,0.0005063219,0.0013950187,0.0019760204,0.0013689877,0.0021282942],"category_scores_gemma":[0.007734466,0.00044265247,0.0016276569,0.004365838,0.000689491,0.0037838216,0.0010963117,0.0015011121,0.0014948738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003064956,0.00025716927,0.0018952854,0.00043859,0.00024490082,0.0001509045,0.0002382496,0.17933176,0.023036078,0.024937108,0.013062911,0.75610054],"study_design_scores_gemma":[0.000025059207,0.000091897266,0.0011655182,0.000030375244,0.00007247851,0.00013738315,0.000028609618,0.9766987,0.004714698,0.011079753,0.0058986135,0.000056908197],"about_ca_topic_score_codex":0.008849991,"about_ca_topic_score_gemma":0.0072598835,"teacher_disagreement_score":0.008849991,"about_ca_system_score_codex":0.0012688785,"about_ca_system_score_gemma":0.0016299821,"threshold_uncertainty_score":0.01759696},"labels":[],"label_agreement":null},{"id":"W2990395788","doi":"10.26615/978-954-452-056-4_119","title":"Self-Attentional Models Application in Task-Oriented Dialogue Generation Systems","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Task (project management); Machine translation; Sequence (biology); Artificial intelligence; Convolution (computer science); Task analysis; Sequence learning; Mechanism (biology); Machine learning; Natural language processing; Human–computer interaction; Artificial neural network","score_opus":0.019001409230749056,"score_gpt":0.22129831982491097,"score_spread":0.2022969105941619,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990395788","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19627964,0.0017994297,0.77929807,0.0006688644,0.00031672255,0.00032507276,0.00056375004,0.013585126,0.007163438],"genre_scores_gemma":[0.83603984,0.00030737137,0.15722577,0.00032009333,0.00008943046,0.0003216176,0.0010971581,0.0004013555,0.0041973717],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99858665,0.00067428855,0.00008238217,0.0003754818,0.00018511628,0.00009614235],"domain_scores_gemma":[0.9972178,0.0016642174,0.00014805583,0.00032040407,0.00050124695,0.00014844589],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027682658,0.00088211993,0.00055564224,0.00080722664,0.00044075915,0.0011577372,0.0009951243,0.001274066,0.0023932524],"category_scores_gemma":[0.006767267,0.0003456031,0.0008224267,0.00044027143,0.0003702362,0.0016534424,0.0012206725,0.0012642153,0.0010524447],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009088467,0.00073615415,0.0060794903,0.00075410696,0.00037011606,0.00042684155,0.002498461,0.37300944,0.050597806,0.009268866,0.011572466,0.5437774],"study_design_scores_gemma":[0.00002431138,0.00013621247,0.0007639359,0.000019316012,0.000040820814,0.000060504706,0.000076058066,0.9851681,0.008109004,0.002923065,0.0026578612,0.00002071399],"about_ca_topic_score_codex":0.004962325,"about_ca_topic_score_gemma":0.0060732234,"teacher_disagreement_score":0.004962325,"about_ca_system_score_codex":0.00073845073,"about_ca_system_score_gemma":0.0008614259,"threshold_uncertainty_score":0.014640152},"labels":[],"label_agreement":null},{"id":"W2990468116","doi":"10.22148/16.051","title":"Annotation Guideline No. 1: Narrative Boundaries Annotation Guide","year":2019,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narrative; Annotation; Natural (archaeology); Cover (algebra); Test (biology); Computer science; World Wide Web; History; Literature; Artificial intelligence; Engineering; Art; Biology","score_opus":0.020478133666557263,"score_gpt":0.2973478058838749,"score_spread":0.2768696722173176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990468116","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044606803,0.0017514606,0.26172838,0.023229644,0.0053543225,0.024257017,0.2400808,0.025800057,0.4133377],"genre_scores_gemma":[0.01735383,0.0024704672,0.42524323,0.012384449,0.00096122414,0.055074725,0.2292656,0.012713205,0.24453318],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9865317,0.004825804,0.0034860526,0.001306993,0.003137187,0.0007123391],"domain_scores_gemma":[0.9084917,0.029958714,0.0029983898,0.009123596,0.04797047,0.0014571062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01712619,0.001425258,0.0013380712,0.0097898785,0.004824449,0.007228232,0.0045004943,0.0064795953,0.15674499],"category_scores_gemma":[0.084727414,0.0020031827,0.001302354,0.0061294236,0.0024725443,0.006576444,0.0061143804,0.0045505376,0.14656083],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000864153,0.000036403802,0.00031864367,0.002163859,0.000008344337,0.0001929357,0.004616411,0.00026503523,0.0022200337,0.01572643,0.93548334,0.03888217],"study_design_scores_gemma":[0.000013569562,0.000005859384,0.00035490026,0.0010694884,0.0000071630852,0.000090636764,0.00081089017,0.00024552038,0.00092120207,0.0044919765,0.9919619,0.000026902859],"about_ca_topic_score_codex":0.022855056,"about_ca_topic_score_gemma":0.041556705,"teacher_disagreement_score":0.15674499,"about_ca_system_score_codex":0.004279104,"about_ca_system_score_gemma":0.013990814,"threshold_uncertainty_score":0.52436423},"labels":[],"label_agreement":null},{"id":"W2990476816","doi":"10.3233/shti190510","title":"AutoScribe: Extracting Clinically Pertinent Information from Patient-Clinician Dialogues","year":2019,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Michael's Hospital; Vector Institute; University of Toronto","funders":"","keywords":"Context (archaeology); Computer science; Medical information; Information retrieval; Natural language processing; Data science; Artificial intelligence; History","score_opus":0.061251424951719605,"score_gpt":0.34707735070678375,"score_spread":0.28582592575506416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990476816","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06457238,0.0051542353,0.82555515,0.002494374,0.0005309613,0.0014653297,0.02399756,0.06614445,0.010085552],"genre_scores_gemma":[0.2044035,0.0018826919,0.75307256,0.0007770539,0.00040232402,0.0007065625,0.030155879,0.002717414,0.0058819484],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969502,0.0012752706,0.0002751725,0.00079180894,0.00060558825,0.00010200724],"domain_scores_gemma":[0.98512876,0.011983573,0.0007534088,0.00081985997,0.0009945396,0.00031988986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038226512,0.0015018,0.0011924385,0.0037885786,0.0008544335,0.0022209361,0.0011739917,0.0015227731,0.007072889],"category_scores_gemma":[0.013989077,0.00078741234,0.0008868268,0.0015945041,0.0005458359,0.0024665007,0.0023730136,0.0011028239,0.0035146414],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012553439,0.0004134122,0.015267863,0.0047056074,0.0003436527,0.0030066413,0.013102308,0.008678473,0.07513986,0.01329589,0.1329534,0.73183745],"study_design_scores_gemma":[0.0004019118,0.0006149001,0.028152926,0.0013763556,0.0005362819,0.0069393916,0.0060215895,0.32437524,0.10606704,0.039152846,0.48589918,0.00046234138],"about_ca_topic_score_codex":0.002337491,"about_ca_topic_score_gemma":0.0038831094,"teacher_disagreement_score":0.007072889,"about_ca_system_score_codex":0.00066849205,"about_ca_system_score_gemma":0.001956815,"threshold_uncertainty_score":0.023661196},"labels":[],"label_agreement":null},{"id":"W2990647120","doi":"10.22148/16.052","title":"Annotation Guideline No. 2: For Annotating Anachronies and Narrative Levels in Fiction","year":2019,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narrative; Annotation; Perspective (graphical); Guideline; Encoding (memory); Computer science; World Wide Web; Literature; Art; Artificial intelligence; Medicine","score_opus":0.03843314680947173,"score_gpt":0.31443756346654733,"score_spread":0.2760044166570756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990647120","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033219505,0.0011201055,0.6950455,0.01010772,0.0054035466,0.010001754,0.07788606,0.023830598,0.14338528],"genre_scores_gemma":[0.07221563,0.0008986213,0.74896365,0.004277873,0.0005349187,0.019028839,0.067364484,0.009570423,0.07714554],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9921482,0.0028291864,0.001822835,0.0011678466,0.0014814491,0.0005504256],"domain_scores_gemma":[0.96218747,0.013566965,0.0013888596,0.0055885203,0.016473107,0.0007952135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009889239,0.0013176347,0.001043153,0.006584985,0.0043273685,0.004271099,0.0017763645,0.0042767706,0.02217668],"category_scores_gemma":[0.037142955,0.0012679574,0.0009323286,0.0038532305,0.0027893255,0.004630474,0.0042234478,0.0034491397,0.021220328],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000455672,0.00018273761,0.0036284397,0.0052381125,0.000053398304,0.0018811277,0.05877521,0.0010789285,0.051754147,0.08705492,0.6340094,0.15588784],"study_design_scores_gemma":[0.000025243453,0.000020548927,0.0018995544,0.0010796518,0.000033640365,0.0003701187,0.0040420447,0.0013161916,0.010455834,0.012564286,0.96811354,0.00007938173],"about_ca_topic_score_codex":0.011593672,"about_ca_topic_score_gemma":0.018715657,"teacher_disagreement_score":0.02217668,"about_ca_system_score_codex":0.0021283126,"about_ca_system_score_gemma":0.005287776,"threshold_uncertainty_score":0.07418841},"labels":[],"label_agreement":null},{"id":"W2990684147","doi":"10.1145/3459104.3459147","title":"Relation Extraction with Synthetic Explanations and Neural Network","year":2021,"lang":"en","type":"article","venue":"2021 International Symposium on Electrical, Electronics and Information Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Relationship extraction; Relation (database); Computer science; Artificial intelligence; Sentence; Artificial neural network; Training set; Set (abstract data type); Machine learning; Natural language processing; Noise (video); Pattern recognition (psychology); Data mining","score_opus":0.0037985368578463768,"score_gpt":0.19099787517394093,"score_spread":0.18719933831609456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2990684147","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27629083,0.0024087855,0.6853288,0.0015670619,0.00037090867,0.00073902076,0.010786249,0.017031403,0.005476886],"genre_scores_gemma":[0.40907866,0.0005408141,0.56165427,0.0002876609,0.000102697944,0.0006399762,0.02461713,0.0003802708,0.002698421],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985128,0.0006502346,0.00010788385,0.00038353156,0.0002883297,0.000057225905],"domain_scores_gemma":[0.99395907,0.004477719,0.00037314734,0.0005575686,0.0005552747,0.000077293626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013907758,0.0013897425,0.00041046197,0.0014129719,0.00041674628,0.0005841163,0.0011670714,0.0011981666,0.0022181906],"category_scores_gemma":[0.0093944045,0.00035424024,0.00084139185,0.0012909399,0.0005123475,0.0012170572,0.00090230134,0.0010018855,0.000747667],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014337227,0.0007086637,0.013948132,0.002951548,0.00040115713,0.0046205786,0.002600397,0.24196495,0.05364848,0.02103014,0.05180894,0.6048833],"study_design_scores_gemma":[0.00016477912,0.00027064822,0.003907126,0.00012516169,0.000116639494,0.0007880464,0.0004733841,0.9142611,0.035557263,0.017435769,0.0268414,0.000058715275],"about_ca_topic_score_codex":0.0026737347,"about_ca_topic_score_gemma":0.005967916,"teacher_disagreement_score":0.0026737347,"about_ca_system_score_codex":0.0006520986,"about_ca_system_score_gemma":0.0006768616,"threshold_uncertainty_score":0.0074205995},"labels":[],"label_agreement":null},{"id":"W2991233348","doi":"10.22148/16.053","title":"Annotating Narrative Levels: Review of Guideline No. 1.","year":2019,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narrative; Annotation; Guideline; Computer science; Linguistics; Political science; Artificial intelligence; Philosophy","score_opus":0.052936451256500996,"score_gpt":0.33169971006662957,"score_spread":0.2787632588101286,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2991233348","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054237386,0.241056,0.13991547,0.43266243,0.058516834,0.028494801,0.013826787,0.0043367986,0.0757673],"genre_scores_gemma":[0.02312615,0.20364398,0.44964206,0.19795099,0.006653211,0.048788734,0.02342168,0.002897043,0.043876156],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9147953,0.0331704,0.02909009,0.0032305107,0.017988412,0.001725356],"domain_scores_gemma":[0.63804066,0.14406471,0.019909432,0.015242056,0.17816591,0.0045771822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11332818,0.0011413101,0.002931807,0.015272484,0.0038671894,0.006877466,0.010492061,0.0089701,0.0060226195],"category_scores_gemma":[0.318194,0.0020595659,0.003961745,0.009295704,0.006464255,0.006946199,0.007238597,0.008390516,0.0080048265],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008527013,0.00007881762,0.00069399044,0.05082839,0.00016353272,0.00041020304,0.007273589,0.0003678853,0.0013454561,0.012996351,0.65564585,0.27011076],"study_design_scores_gemma":[0.000051689,0.00004513609,0.00093306816,0.102310814,0.00026241003,0.00035959127,0.0015335868,0.00017331,0.0009510427,0.0049476046,0.88834906,0.00008259736],"about_ca_topic_score_codex":0.03107907,"about_ca_topic_score_gemma":0.07949231,"teacher_disagreement_score":0.11332818,"about_ca_system_score_codex":0.00994156,"about_ca_system_score_gemma":0.06419393,"threshold_uncertainty_score":0.5993439},"labels":[],"label_agreement":null},{"id":"W2991485158","doi":"10.48550/arxiv.1911.13280","title":"Deconstructing and reconstructing word embedding algorithms","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Word2vec; Pointwise mutual information; Word (group theory); Word embedding; Computer science; Pointwise; Embedding; Algorithm; Feature (linguistics); Construct (python library); Artificial intelligence; Natural language processing; Mutual information; Mathematics; Linguistics","score_opus":0.07496298646145166,"score_gpt":0.2028099849556326,"score_spread":0.12784699849418094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2991485158","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03179516,0.0003209357,0.9656257,0.00029527294,0.00004607139,0.000042331747,0.0001262152,0.0010058555,0.00074251293],"genre_scores_gemma":[0.35372865,0.0004844667,0.639541,0.0002339524,0.00011395482,0.00023649905,0.0015999748,0.00050751196,0.0035541267],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9977064,0.0010311627,0.00019178324,0.0005683743,0.000370564,0.00013170866],"domain_scores_gemma":[0.9947848,0.0024134004,0.00031154763,0.0014930117,0.0008628471,0.00013440553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032189097,0.0013899674,0.00093751913,0.0012270887,0.00050844747,0.0017850447,0.0013335694,0.0015610015,0.0014764126],"category_scores_gemma":[0.018022558,0.00053000817,0.00081763975,0.0010751135,0.001532525,0.0052348226,0.0032183407,0.0027891994,0.0017238698],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032283308,0.00021973981,0.0046796952,0.00031862254,0.0001264528,0.00013207148,0.00067031593,0.28049418,0.022062264,0.08380237,0.005277051,0.6018943],"study_design_scores_gemma":[0.000028328192,0.00010505755,0.00042195,0.000026011876,0.000013258253,0.00009905877,0.00011877478,0.9211272,0.013687043,0.061145518,0.0032008858,0.000026794796],"about_ca_topic_score_codex":0.0009767774,"about_ca_topic_score_gemma":0.0018465399,"teacher_disagreement_score":0.0032189097,"about_ca_system_score_codex":0.00063217664,"about_ca_system_score_gemma":0.0010709057,"threshold_uncertainty_score":0.017023385},"labels":[],"label_agreement":null},{"id":"W2992721737","doi":"10.48550/arxiv.1912.01706","title":"A Robust Self-Learning Method for Fully Unsupervised Cross-Lingual Mappings of Word Embeddings: Making the Method Robustly Reproducible as Well","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Word (group theory); Computer science; Artificial intelligence; Unsupervised learning; Natural language processing; Mathematics; Geometry","score_opus":0.107240782913354,"score_gpt":0.2717381276094857,"score_spread":0.1644973446961317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2992721737","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012879772,0.00017982999,0.9829816,0.0002591968,0.00014706957,0.00010927053,0.00019529242,0.0020508117,0.0011970605],"genre_scores_gemma":[0.27884907,0.00016188635,0.7107323,0.00051022135,0.00020981059,0.0006886173,0.0019834307,0.0021349597,0.004729761],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98936635,0.005162349,0.0005351593,0.003348801,0.001323631,0.00026367174],"domain_scores_gemma":[0.9731415,0.008769376,0.00091421435,0.013373193,0.0034096679,0.00039195275],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01121361,0.0015774115,0.0013553746,0.0013436858,0.000986888,0.002945716,0.0031790046,0.0027854445,0.0038237388],"category_scores_gemma":[0.048980106,0.00089150685,0.002098112,0.001404562,0.0018233992,0.0049204347,0.0047830627,0.0042229644,0.004749131],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069896755,0.00058209064,0.0075794253,0.00055191526,0.0011855668,0.00029743978,0.0012172437,0.2408698,0.03480838,0.056458887,0.018118462,0.6376318],"study_design_scores_gemma":[0.00010410238,0.0001514425,0.0009793695,0.000041294894,0.000071690985,0.00019606041,0.00012729323,0.9278603,0.01746353,0.046452675,0.006471485,0.00008078189],"about_ca_topic_score_codex":0.0019003592,"about_ca_topic_score_gemma":0.0023482868,"teacher_disagreement_score":0.9887864,"about_ca_system_score_codex":0.0007731679,"about_ca_system_score_gemma":0.0018325966,"threshold_uncertainty_score":0.059304},"labels":[],"label_agreement":null},{"id":"W2993877978","doi":"10.22148/16.057","title":"Annotating Narrative Levels: Review of Guideline No. 4","year":2019,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narrative; Annotation; Perspective (graphical); Guideline; Computer science; Narrative review; World Wide Web; Psychology; Linguistics; Political science; Artificial intelligence; Philosophy","score_opus":0.05391653072855011,"score_gpt":0.33259020478020473,"score_spread":0.2786736740516546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2993877978","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00674624,0.13557027,0.20721272,0.4359478,0.10471696,0.009261688,0.010084802,0.0055559725,0.084903575],"genre_scores_gemma":[0.03554369,0.12634468,0.47955948,0.23110284,0.013538265,0.021821778,0.02586683,0.0058872504,0.06033522],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9331995,0.02445969,0.016949616,0.0037889616,0.01965037,0.0019517831],"domain_scores_gemma":[0.68457264,0.09144962,0.012636404,0.0151043115,0.1919891,0.0042478987],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10276916,0.0009979474,0.0021175514,0.013515101,0.0045428975,0.009136911,0.008900743,0.008685789,0.00590626],"category_scores_gemma":[0.2758678,0.0016163999,0.0030372096,0.009206225,0.0063182283,0.006945492,0.007473787,0.009881469,0.0077988417],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007094098,0.00006401726,0.0007092884,0.017735956,0.00009051554,0.00033480712,0.0067832693,0.00030914112,0.0023710788,0.02595146,0.75983185,0.18574768],"study_design_scores_gemma":[0.000022031092,0.000029635925,0.00079856155,0.023446633,0.000109613466,0.00018264182,0.0012510505,0.000170714,0.0010916817,0.004909776,0.9679281,0.000059522707],"about_ca_topic_score_codex":0.03359148,"about_ca_topic_score_gemma":0.061958976,"teacher_disagreement_score":0.10276916,"about_ca_system_score_codex":0.0089656,"about_ca_system_score_gemma":0.04490562,"threshold_uncertainty_score":0.54350173},"labels":[],"label_agreement":null},{"id":"W2994747004","doi":"10.1016/j.elerap.2019.100917","title":"A neural graph embedding approach for selecting review sentences","year":2019,"lang":"en","type":"article","venue":"Electronic Commerce Research and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Embedding; Artificial intelligence; Graph; Natural language processing; Machine learning; Theoretical computer science","score_opus":0.06821664792207113,"score_gpt":0.3889785058434018,"score_spread":0.3207618579213307,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2994747004","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.066358805,0.008711264,0.89885265,0.002247384,0.0010071697,0.0009766674,0.009244142,0.0058004125,0.006801428],"genre_scores_gemma":[0.4659554,0.0029860758,0.49368846,0.0006163544,0.0012194824,0.00094337034,0.021109026,0.00044601219,0.013035901],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985594,0.00051230093,0.00017392228,0.00037951814,0.0002886736,0.00008616446],"domain_scores_gemma":[0.99603695,0.0023309998,0.00026629816,0.00026951171,0.00097314036,0.00012305818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016559401,0.0011057805,0.0010485796,0.006607589,0.0006672496,0.0013400989,0.0011402023,0.0015158079,0.0032275456],"category_scores_gemma":[0.006696054,0.0003670809,0.0011154622,0.0049879197,0.0002996875,0.002311104,0.0010270668,0.0012617512,0.0018351027],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074624334,0.0006502332,0.005712606,0.0011253918,0.0005785183,0.00032850273,0.00033884583,0.018965188,0.023357615,0.00912013,0.0605557,0.878521],"study_design_scores_gemma":[0.00015486582,0.00034203636,0.006513271,0.0001801136,0.0005942777,0.00044942083,0.00024672953,0.9356139,0.009158153,0.02409374,0.022567317,0.00008624868],"about_ca_topic_score_codex":0.0052329535,"about_ca_topic_score_gemma":0.014811214,"teacher_disagreement_score":0.006607589,"about_ca_system_score_codex":0.00066047045,"about_ca_system_score_gemma":0.0015701604,"threshold_uncertainty_score":0.010797143},"labels":[],"label_agreement":null},{"id":"W2995652189","doi":"10.1145/3458553.3458560","title":"Report on the First HIPstIR Workshop on the Future of Information Retrieval","year":2019,"lang":"en","type":"preprint","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Mainstream; Library science; Computer science; Information retrieval; Political science","score_opus":0.02401280320501431,"score_gpt":0.25147593058079315,"score_spread":0.22746312737577884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2995652189","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025318895,0.05978018,0.099225305,0.2739617,0.26206362,0.0033131125,0.019349864,0.003778427,0.25320888],"genre_scores_gemma":[0.0398601,0.018836938,0.035453454,0.03228858,0.044709902,0.001219163,0.025462646,0.0021316474,0.80003756],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953661,0.0014315195,0.0001636433,0.00052489695,0.0018984557,0.0006153552],"domain_scores_gemma":[0.98550457,0.003851751,0.0004790234,0.0011717387,0.0044520716,0.0045408965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011935525,0.0011724806,0.0010070816,0.0018850467,0.0019885246,0.008116216,0.0021087846,0.0036639227,0.0740065],"category_scores_gemma":[0.014494858,0.000508943,0.0016652762,0.0017183954,0.0008951861,0.007167616,0.00618181,0.0058891675,0.041534163],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017027013,0.00020677385,0.00035891053,0.00026464165,0.000033084605,0.00016485906,0.00048145806,0.00032809222,0.0011469796,0.0045193667,0.9355482,0.05677735],"study_design_scores_gemma":[0.000043621687,0.000112324764,0.00089779805,0.0001800077,0.000026822105,0.00008207657,0.00049224624,0.00045893952,0.0007330881,0.0032741255,0.9936621,0.000036791134],"about_ca_topic_score_codex":0.006561175,"about_ca_topic_score_gemma":0.009988232,"teacher_disagreement_score":0.0740065,"about_ca_system_score_codex":0.001716596,"about_ca_system_score_gemma":0.0043258113,"threshold_uncertainty_score":0.24757642},"labels":[],"label_agreement":null},{"id":"W2996068536","doi":"","title":"Language GANs Falling Short","year":2020,"lang":"en","type":"article","venue":"International Conference on Learning Representations","topic":"Topic Modeling","field":"Computer Science","cited_by":89,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; Université de Montréal; McGill University","funders":"","keywords":"Softmax function; Computer science; Sample (material); Inference; Metric (unit); Natural language generation; Quality (philosophy); Artificial intelligence; Diversity (politics); Machine learning; Ground truth; Fallacy; Sampling bias; Sample size determination; Statistics; Mathematics; Artificial neural network; Natural language","score_opus":0.09059197135597538,"score_gpt":0.3557337931421951,"score_spread":0.2651418217862197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996068536","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023175271,0.0018397301,0.95188665,0.001507418,0.00041469143,0.00007170703,0.00070303096,0.0038049794,0.016596366],"genre_scores_gemma":[0.7756114,0.0020342558,0.1863676,0.0025856188,0.0004952698,0.0003153138,0.0028722193,0.0016697226,0.028048638],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991036,0.00031588407,0.000036878235,0.00024037359,0.00020204937,0.00010123044],"domain_scores_gemma":[0.9978853,0.001295704,0.000103420025,0.0004374586,0.00019764212,0.00008061722],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013926069,0.0009105079,0.00081318367,0.00039708003,0.00029860312,0.001317119,0.0011722283,0.0009872721,0.007294648],"category_scores_gemma":[0.006190041,0.00049282104,0.00074789044,0.00036660617,0.0009094697,0.002724502,0.0017490134,0.0028023713,0.0026293034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002772564,0.000084975596,0.0016721712,0.00025350985,0.00014729406,0.00023379044,0.00022417582,0.52835166,0.009202289,0.20096654,0.026598502,0.23198786],"study_design_scores_gemma":[0.000019018384,0.00005118941,0.00033382775,0.00003230915,0.000015845359,0.00011710469,0.00002176469,0.89556485,0.002078778,0.09015306,0.011597142,0.000015207294],"about_ca_topic_score_codex":0.0016566851,"about_ca_topic_score_gemma":0.0031437834,"teacher_disagreement_score":0.007294648,"about_ca_system_score_codex":0.0008162557,"about_ca_system_score_gemma":0.00082641625,"threshold_uncertainty_score":0.024403036},"labels":[],"label_agreement":null},{"id":"W2996350961","doi":"10.1186/s12911-019-0980-z","title":"Improving clinical named entity recognition in Chinese using the graphical and phonetic feature","year":2019,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"University of Manchester","keywords":"Computer science; Natural language processing; Artificial intelligence; Feature (linguistics); Named-entity recognition; Pinyin; Chinese characters; Embedding; Information retrieval; Linguistics","score_opus":0.04025595392247056,"score_gpt":0.3433896400771219,"score_spread":0.30313368615465136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996350961","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5482379,0.0031774677,0.42396775,0.0028673643,0.00041965317,0.00046061102,0.00733759,0.005293926,0.008237657],"genre_scores_gemma":[0.87177765,0.001079608,0.11586983,0.00018804516,0.00010630704,0.00012158581,0.008316227,0.00012613948,0.0024145925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978886,0.0007012654,0.0002769756,0.00063651317,0.00035182832,0.00014474163],"domain_scores_gemma":[0.9942992,0.0034323228,0.00041614284,0.0006176402,0.0011107275,0.00012392887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036879561,0.0010008786,0.0006942839,0.0027350555,0.00057416415,0.0014858001,0.00077426864,0.00056807674,0.0020257372],"category_scores_gemma":[0.011080037,0.00019765674,0.0012321062,0.0030988527,0.00038467292,0.0033481726,0.0013032977,0.0007291127,0.0012224783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071358326,0.00021638916,0.089619264,0.0007409358,0.00040442977,0.0008586481,0.001006221,0.045180947,0.01566616,0.0038744311,0.011175405,0.8305436],"study_design_scores_gemma":[0.00007774349,0.00027698008,0.08540117,0.0001482412,0.0009182552,0.0009856396,0.00086466304,0.8286107,0.052201685,0.009026212,0.021288056,0.00020060057],"about_ca_topic_score_codex":0.016130442,"about_ca_topic_score_gemma":0.011148192,"teacher_disagreement_score":0.016130442,"about_ca_system_score_codex":0.0008974368,"about_ca_system_score_gemma":0.0015512443,"threshold_uncertainty_score":0.03207314},"labels":[],"label_agreement":null},{"id":"W2996912688","doi":"","title":"Discreteness in Neural Natural Language Processing","year":2019,"lang":"en","type":"article","venue":"Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Artificial neural network; Focus (optics); Process (computing); Point (geometry); Space (punctuation); Natural language processing; Natural language; Machine learning; Programming language; Mathematics","score_opus":0.03199366781166351,"score_gpt":0.4167666648137828,"score_spread":0.38477299700211925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996912688","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036146885,0.025255864,0.94389904,0.004830928,0.0004760487,0.000049627157,0.0003468149,0.00029857448,0.021228457],"genre_scores_gemma":[0.31938472,0.06627134,0.57834977,0.0033230092,0.0039652125,0.00079997577,0.0014531821,0.0004947727,0.025958002],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986045,0.0005509247,0.0001267985,0.0002790199,0.00037881333,0.00006007013],"domain_scores_gemma":[0.99672264,0.0027253984,0.00010360211,0.00022887686,0.00016484817,0.00005470367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026968534,0.0006411536,0.0008731439,0.0014517121,0.00051706936,0.0030947884,0.0012737924,0.001437101,0.0056475857],"category_scores_gemma":[0.008808486,0.00055805,0.00082591287,0.002033613,0.0034380779,0.007027759,0.0016777097,0.0045277374,0.0011701041],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008682474,0.000012571715,0.00017173054,0.00018680556,0.00001803061,0.00004403601,0.00011964623,0.008050798,0.00029567437,0.9580468,0.0028933098,0.030151831],"study_design_scores_gemma":[0.0000022906881,0.0000063644793,0.00010621962,0.00006262793,0.0000036876781,0.000046745794,0.000017196375,0.031842202,0.00017938491,0.9566795,0.011044798,0.000008947861],"about_ca_topic_score_codex":0.0017103464,"about_ca_topic_score_gemma":0.0014868022,"teacher_disagreement_score":0.0056475857,"about_ca_system_score_codex":0.0017098791,"about_ca_system_score_gemma":0.0009223366,"threshold_uncertainty_score":0.018893123},"labels":[],"label_agreement":null},{"id":"W2996917304","doi":"10.1609/aaai.v34i05.6423","title":"Relation Extraction with Convolutional Network over Learnable Syntax-Transport Graph","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Beijing Advanced Innovation Center for Big Data and Brain Computing; Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China","keywords":"Computer science; Graph; Dependency graph; Dependency (UML); Artificial intelligence; Relationship extraction; Theoretical computer science; Abstract syntax; Natural language processing; Syntax; Information extraction","score_opus":0.07050531782830213,"score_gpt":0.26828117828860537,"score_spread":0.19777586046030324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996917304","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06186781,0.0019225599,0.8962816,0.00077382533,0.00021665628,0.0002517755,0.004062501,0.026798416,0.007824976],"genre_scores_gemma":[0.49622864,0.0016851814,0.46596295,0.00064371416,0.00012385375,0.00029183723,0.019576387,0.000841808,0.014645564],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99946207,0.000047528818,0.000033772445,0.00026518598,0.00011115753,0.00008032127],"domain_scores_gemma":[0.99953663,0.00014248733,0.000064269225,0.00013101309,0.000094492796,0.000031140873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00043424888,0.0019585832,0.00093168346,0.003032214,0.00062927086,0.0010600247,0.0016888124,0.0013405291,0.0042993105],"category_scores_gemma":[0.0014070785,0.00068604294,0.0019286247,0.003328356,0.00069090025,0.0036466226,0.0014456525,0.0019078047,0.0027080574],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000302578,0.00024998817,0.003105888,0.00045074147,0.00022093824,0.0005663618,0.00024389275,0.07678773,0.029224446,0.021769159,0.029368835,0.8377094],"study_design_scores_gemma":[0.000028590737,0.000067854286,0.0019509066,0.00006125425,0.00013783918,0.00021818321,0.00006880882,0.923698,0.01609429,0.04796935,0.009666414,0.000038437727],"about_ca_topic_score_codex":0.01891886,"about_ca_topic_score_gemma":0.030481648,"teacher_disagreement_score":0.01891886,"about_ca_system_score_codex":0.0015054666,"about_ca_system_score_gemma":0.0016355569,"threshold_uncertainty_score":0.037617505},"labels":[],"label_agreement":null},{"id":"W2997103344","doi":"10.48550/arxiv.1912.13082","title":"The Shmoop Corpus: A Dataset of Stories with Loosely Aligned Summaries","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Automatic summarization; Paragraph; Natural language processing; Artificial intelligence; Set (abstract data type); Reading comprehension; Construct (python library); Exploit; Reading (process); Comprehension; Sentence; Exposition (narrative); Linguistics; World Wide Web; Literature; Art; Programming language","score_opus":0.06948629363307242,"score_gpt":0.18629075462889852,"score_spread":0.1168044609958261,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997103344","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12169425,0.004690554,0.01503675,0.0021376363,0.0005311516,0.0008962261,0.82516336,0.007425357,0.022424636],"genre_scores_gemma":[0.05919181,0.0006554689,0.022045843,0.00019891754,0.00012680715,0.0011500048,0.9099254,0.0004675693,0.006238087],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99894696,0.00035619715,0.00011205633,0.00025318837,0.00025746855,0.00007410611],"domain_scores_gemma":[0.9966569,0.0018441536,0.00021090632,0.00045097005,0.00061730004,0.00021970332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007547138,0.0012700604,0.00056436664,0.0034633495,0.0013449033,0.0010943662,0.0014107581,0.0018824098,0.014316042],"category_scores_gemma":[0.008480022,0.00033186344,0.00068611326,0.003907725,0.000761554,0.0018202298,0.0020269984,0.0012638561,0.008592331],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000693859,0.00047107908,0.007552605,0.005149332,0.00016713016,0.0018378265,0.0037261033,0.0031472552,0.008806587,0.0048040124,0.81363165,0.15001264],"study_design_scores_gemma":[0.00043112724,0.0002484925,0.04176983,0.0005297911,0.00011639692,0.0018139177,0.004015024,0.013715132,0.010220663,0.007465048,0.919533,0.00014149443],"about_ca_topic_score_codex":0.009326194,"about_ca_topic_score_gemma":0.022967828,"teacher_disagreement_score":0.014316042,"about_ca_system_score_codex":0.0007182152,"about_ca_system_score_gemma":0.0012031256,"threshold_uncertainty_score":0.047891974},"labels":[],"label_agreement":null},{"id":"W2997152227","doi":"10.1111/jcal.12531","title":"Automatic identification of knowledge‐transforming content in argument essays developed from multiple sources","year":2021,"lang":"en","type":"article","venue":"Journal of Computer Assisted Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kwantlen Polytechnic University; Simon Fraser University","funders":"Simon Fraser University","keywords":"Argumentative; Paraphrase; Computer science; Domain knowledge; Argument (complex analysis); Identification (biology); Natural language processing; Typology; Linguistics; Artificial intelligence; Psychology; Sociology","score_opus":0.045843568205954796,"score_gpt":0.2673284600863535,"score_spread":0.22148489188039872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997152227","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89991087,0.000420705,0.0927494,0.00034896456,0.00008711036,0.00029923793,0.001044665,0.0016412694,0.0034977635],"genre_scores_gemma":[0.9018342,0.00012019356,0.09581057,0.000031056326,0.000033647306,0.00011610209,0.0013079436,0.000091561655,0.0006546878],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983165,0.00077338563,0.00019628032,0.00041020178,0.0002403767,0.00006331914],"domain_scores_gemma":[0.9654497,0.027125636,0.00267509,0.0012030009,0.003186813,0.00035972992],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033151964,0.0004824292,0.0003537875,0.0039099753,0.0005666318,0.0024770764,0.0005668141,0.0008898388,0.0017873072],"category_scores_gemma":[0.026288494,0.0002313606,0.0003628992,0.0013249695,0.0004223154,0.002144853,0.001149501,0.00087791146,0.0008150872],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009066973,0.0004915837,0.09661598,0.0010229722,0.00013333347,0.0011587365,0.008843352,0.004663278,0.09901355,0.0026872791,0.0037910664,0.78067225],"study_design_scores_gemma":[0.0001641535,0.000634946,0.2130687,0.0006119869,0.00035552966,0.0019945481,0.010261475,0.58486223,0.15437464,0.013006261,0.020506363,0.00015911699],"about_ca_topic_score_codex":0.00040075777,"about_ca_topic_score_gemma":0.000748203,"teacher_disagreement_score":0.0039099753,"about_ca_system_score_codex":0.00044284321,"about_ca_system_score_gemma":0.000544088,"threshold_uncertainty_score":0.017532647},"labels":[],"label_agreement":null},{"id":"W2997327944","doi":"10.48550/arxiv.1912.12481","title":"Robust Cross-lingual Embeddings from Parallel Sentences","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Sentence; Word (group theory); Embedding; Inference; Speech recognition; Mathematics","score_opus":0.12005334569043298,"score_gpt":0.21305352125842036,"score_spread":0.09300017556798738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997327944","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07982123,0.0026384825,0.8980313,0.00050073443,0.0006512048,0.00019040005,0.0029292032,0.010158254,0.005079226],"genre_scores_gemma":[0.5507572,0.0023157345,0.39570943,0.0004941635,0.0006938744,0.00043318895,0.02783574,0.0023292673,0.019431371],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99839896,0.00045124762,0.00014291584,0.0006379142,0.00022734927,0.00014154856],"domain_scores_gemma":[0.99793136,0.0005580076,0.0001780332,0.00069401274,0.0005660673,0.00007256898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016039723,0.002333054,0.001225586,0.0025622037,0.0005862078,0.0016145125,0.0014593522,0.001150495,0.0054577603],"category_scores_gemma":[0.00605435,0.0006706914,0.0015134289,0.0025295937,0.0005685699,0.0050745066,0.0035906467,0.0019642399,0.008707959],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046338275,0.000360554,0.0028671715,0.00057866774,0.0003503689,0.0003532961,0.000672783,0.02784861,0.037774835,0.009492246,0.027346013,0.8918919],"study_design_scores_gemma":[0.00009843968,0.00048602404,0.0055151167,0.00015628875,0.0003061208,0.0008711793,0.001156098,0.8389808,0.04743766,0.06100063,0.043817,0.00017463743],"about_ca_topic_score_codex":0.0024631545,"about_ca_topic_score_gemma":0.004589908,"teacher_disagreement_score":0.0054577603,"about_ca_system_score_codex":0.00050573004,"about_ca_system_score_gemma":0.0012127403,"threshold_uncertainty_score":0.018258035},"labels":[],"label_agreement":null},{"id":"W2997394441","doi":"10.1609/aaai.v34i05.6292","title":"P-SIF: Document Embeddings Using Partition Averaging","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute for Catastrophic Loss Reduction; Microsoft Research","keywords":"Computer science; Simple (philosophy); Word (group theory); Set (abstract data type); Correctness; Representation (politics); Partition (number theory); Natural language processing; Artificial intelligence; Algorithm; Pattern recognition (psychology); Mathematics; Combinatorics","score_opus":0.13534460733668546,"score_gpt":0.31113987685313105,"score_spread":0.17579526951644558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997394441","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016257178,0.0010317999,0.97376627,0.00022492159,0.00018133798,0.00019056907,0.001099742,0.005908148,0.0013400207],"genre_scores_gemma":[0.21038893,0.0012341986,0.7725758,0.00031927472,0.00035733895,0.00071462296,0.007361595,0.0008612505,0.006187007],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989778,0.00023224732,0.00009210547,0.00034391915,0.0002708142,0.00008326107],"domain_scores_gemma":[0.99824274,0.00065247226,0.00015616919,0.0004695818,0.0004089185,0.00007017555],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015753628,0.0018410209,0.0011908562,0.0029356298,0.00058231305,0.0012209748,0.0018495367,0.0012963895,0.003395802],"category_scores_gemma":[0.0069595035,0.00045714766,0.0013273724,0.0033930873,0.00054094405,0.004544405,0.00135818,0.0018241588,0.002595276],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018271622,0.00015728157,0.0018535929,0.00021090018,0.00017599638,0.000104577055,0.00025746654,0.04725923,0.0070631467,0.009052941,0.020771893,0.9129102],"study_design_scores_gemma":[0.00007195415,0.00023334542,0.0012638269,0.000046992478,0.00008394821,0.00035207474,0.00011736301,0.93507874,0.009928769,0.037459668,0.015288122,0.00007528917],"about_ca_topic_score_codex":0.0065984745,"about_ca_topic_score_gemma":0.008332763,"teacher_disagreement_score":0.0065984745,"about_ca_system_score_codex":0.00087632687,"about_ca_system_score_gemma":0.0015132315,"threshold_uncertainty_score":0.013120115},"labels":[],"label_agreement":null},{"id":"W2997405211","doi":"10.1609/aaai.v34i05.6230","title":"Modelling Sentence Pairs via Reinforcement Learning: An Actor-Critic Approach to Learn the Irrelevant Words","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Reinforcement learning; Sentence; Computer science; Artificial intelligence; Natural language processing; Task (project management); Representation (politics); Paraphrase; Inference; Parsing; Similarity (geometry); Machine learning","score_opus":0.13772920933430416,"score_gpt":0.287427590337229,"score_spread":0.14969838100292482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997405211","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030381924,0.0002902304,0.96638966,0.00046313848,0.00008443658,0.00010507621,0.0000542772,0.0007606376,0.0014707019],"genre_scores_gemma":[0.83648884,0.0001743033,0.15898864,0.0004228963,0.000081908685,0.00037456086,0.00018949047,0.00014032911,0.0031390607],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991308,0.00043887348,0.000038731945,0.000205171,0.00011776914,0.00006869898],"domain_scores_gemma":[0.9973732,0.0018780223,0.00021451717,0.00014599673,0.00025587325,0.00013235872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024464026,0.0012616052,0.0014246005,0.00045328206,0.0004073519,0.00082408317,0.0021486816,0.0015238523,0.0020860257],"category_scores_gemma":[0.0077851946,0.0006464272,0.0006454342,0.00041799323,0.0010354379,0.0017240382,0.0011110281,0.0024475805,0.00046708595],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024983336,0.00014809099,0.0011282453,0.00008933666,0.00010282959,0.00018076593,0.00016754091,0.92778933,0.0040708077,0.0097498335,0.0015751483,0.054748323],"study_design_scores_gemma":[0.000009810585,0.000018054754,0.000031795542,0.0000022968936,0.000006263816,0.0000068790555,0.0000032621724,0.9967424,0.00028310667,0.0027767606,0.000115795825,0.0000035511832],"about_ca_topic_score_codex":0.0036829154,"about_ca_topic_score_gemma":0.0040882947,"teacher_disagreement_score":0.0036829154,"about_ca_system_score_codex":0.0009783603,"about_ca_system_score_gemma":0.0011909016,"threshold_uncertainty_score":0.012937963},"labels":[],"label_agreement":null},{"id":"W2997520214","doi":"10.1007/978-981-15-1956-7_15","title":"A Conditional VAE-Based Conversation Model","year":2019,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Conversation; Computer science; Sequence (biology); Artificial intelligence; Block (permutation group theory); Latent variable; Domain (mathematical analysis); Encoder; Autoencoder; Natural language processing; Machine learning; Artificial neural network; Psychology; Mathematics; Communication","score_opus":0.056409538974410145,"score_gpt":0.2836610419750884,"score_spread":0.22725150300067826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997520214","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00863943,0.00047553497,0.97826004,0.0008754592,0.00024658878,0.00012230698,0.0017327077,0.0016769428,0.007971029],"genre_scores_gemma":[0.56446445,0.0010551986,0.38670757,0.000654821,0.0005534666,0.0008237291,0.006265603,0.00080846803,0.038666744],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975611,0.0011534179,0.00011778591,0.0005936579,0.00035356352,0.00022056424],"domain_scores_gemma":[0.9950793,0.0035718733,0.00011711643,0.00042184215,0.00065039156,0.00015953266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029313136,0.0007659176,0.0013731917,0.0013802495,0.0012005938,0.002755529,0.003319183,0.0021174448,0.011881207],"category_scores_gemma":[0.0101841735,0.0008910351,0.001706399,0.0018553956,0.00081687426,0.004783428,0.0022627423,0.0026676902,0.0058829593],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092825026,0.0002845862,0.0018850547,0.0004427102,0.00032675362,0.00029385046,0.0013321576,0.38415846,0.0055757784,0.35205382,0.02771857,0.22500002],"study_design_scores_gemma":[0.000016069798,0.000020626889,0.00014765722,0.000018986513,0.000036777754,0.000063437634,0.000037209255,0.95319957,0.0005233735,0.042006973,0.003906615,0.000022710447],"about_ca_topic_score_codex":0.011859416,"about_ca_topic_score_gemma":0.0102224145,"teacher_disagreement_score":0.011881207,"about_ca_system_score_codex":0.001394658,"about_ca_system_score_gemma":0.0019054344,"threshold_uncertainty_score":0.039746583},"labels":[],"label_agreement":null},{"id":"W2997726005","doi":"10.1007/s00521-019-04681-0","title":"Active neural learners for text with dual supervision","year":2020,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Compute Canada; Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada; Boeing","keywords":"Computer science; Artificial intelligence; Task (project management); Class (philosophy); Convolutional neural network; Machine learning; Word (group theory); Feature (linguistics); Dual (grammatical number); Natural language processing; Recall; Artificial neural network; Information retrieval","score_opus":0.03664262998426257,"score_gpt":0.27709543575545115,"score_spread":0.24045280577118858,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997726005","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02908455,0.0025142115,0.95903456,0.0014694077,0.00053746265,0.00014853814,0.0011245364,0.0031548978,0.0029318994],"genre_scores_gemma":[0.62414056,0.0015556678,0.33712208,0.0006422224,0.0017677066,0.00059770496,0.005816575,0.00082672556,0.027530806],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99861586,0.00042390582,0.000113719885,0.00045985432,0.0002545977,0.00013205204],"domain_scores_gemma":[0.9932829,0.0046390025,0.0002545941,0.0007343682,0.0008465522,0.000242633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031850336,0.0013671851,0.001576549,0.0017208715,0.0010212084,0.0019947353,0.00312489,0.0035361466,0.0064782137],"category_scores_gemma":[0.011476049,0.00082669355,0.0012850894,0.0016959646,0.00080052746,0.0063914144,0.002931354,0.0047918116,0.0033934596],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013133405,0.0008433064,0.0018606081,0.00077379314,0.00027496237,0.0002453754,0.00042820032,0.103395075,0.011861779,0.026866019,0.032286063,0.81985146],"study_design_scores_gemma":[0.00004151846,0.000061697945,0.0002276004,0.00003386283,0.000045590932,0.00004038501,0.00003634625,0.96424395,0.003133546,0.029590715,0.0025309862,0.000013725143],"about_ca_topic_score_codex":0.0029471251,"about_ca_topic_score_gemma":0.006045694,"teacher_disagreement_score":0.0064782137,"about_ca_system_score_codex":0.0009067087,"about_ca_system_score_gemma":0.0013281286,"threshold_uncertainty_score":0.021671832},"labels":[],"label_agreement":null},{"id":"W2997919746","doi":"10.1609/aaai.v34i05.6304","title":"Leveraging Multi-Token Entities in Document-Level Named Entity Recognition","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Fundamental Research Funds for the Central Universities; Renmin University of China; National Natural Science Foundation of China","keywords":"Security token; Computer science; Context (archaeology); Sentence; Task (project management); Natural language processing; Artificial intelligence; Named-entity recognition; Relevance (law); Annotation; Process (computing); Entity linking; Information retrieval; Knowledge base; Programming language","score_opus":0.24879956231443004,"score_gpt":0.30830610231452815,"score_spread":0.05950654000009811,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997919746","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041372433,0.0029094392,0.9400608,0.000539826,0.00034833152,0.00021903636,0.0019035194,0.008688387,0.0039581917],"genre_scores_gemma":[0.4930029,0.0018202197,0.4846482,0.00053313305,0.00031782893,0.00016694427,0.010121705,0.0005129897,0.008876036],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986469,0.00026372308,0.00014675054,0.0006271023,0.00020413875,0.0001113423],"domain_scores_gemma":[0.9975339,0.0010878097,0.00027048288,0.0005628554,0.00045072768,0.00009423119],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022196486,0.0011740078,0.0012292495,0.0031392956,0.0006760248,0.0014704682,0.0018464881,0.001499543,0.0020403417],"category_scores_gemma":[0.0042543085,0.00041431276,0.001341505,0.002951437,0.00063670654,0.007075618,0.0017469566,0.0017804353,0.0026023681],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078731484,0.00034713122,0.008835184,0.0007454527,0.00036645308,0.0013666953,0.00059293234,0.039807312,0.061773337,0.01030858,0.020108955,0.8549606],"study_design_scores_gemma":[0.00005103833,0.00030340411,0.009817806,0.00014102674,0.0005723107,0.0018259123,0.0003064707,0.82395625,0.08247399,0.030409893,0.04992615,0.00021584077],"about_ca_topic_score_codex":0.0041031837,"about_ca_topic_score_gemma":0.008044108,"teacher_disagreement_score":0.0041031837,"about_ca_system_score_codex":0.0006700845,"about_ca_system_score_gemma":0.0006846286,"threshold_uncertainty_score":0.011738777},"labels":[],"label_agreement":null},{"id":"W2998085019","doi":"10.1609/aaai.v34i05.6308","title":"Bayes-Adaptive Monte-Carlo Planning and Learning for Goal-Oriented Dialogues","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Iran Telecommunication Research Center; Ministry of Science and ICT, South Korea; Ministry of Trade, Industry and Energy; National Research Foundation","keywords":"Reinforcement learning; Computer science; Artificial intelligence; Machine learning; Task (project management); Negotiation; Bayesian probability; Bayes' theorem; Domain (mathematical analysis); Mathematics","score_opus":0.1257072937796065,"score_gpt":0.29278048377991567,"score_spread":0.16707319000030918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2998085019","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014175839,0.00015151063,0.98300606,0.00020660264,0.000029053552,0.0000527234,0.000025861867,0.00033772708,0.0020146125],"genre_scores_gemma":[0.6508335,0.00017337823,0.34600642,0.0002014359,0.00005187273,0.00025751945,0.00012395722,0.00012453094,0.0022273639],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986143,0.00071887055,0.000064164626,0.00021791672,0.00027581613,0.00010885228],"domain_scores_gemma":[0.99534,0.0037321632,0.00023423,0.00020986052,0.000319283,0.0001644607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002749465,0.00069755065,0.000947461,0.000502684,0.00053699734,0.00081290246,0.0015164027,0.0012171709,0.0025049734],"category_scores_gemma":[0.011400778,0.00064784574,0.00057057315,0.0003881044,0.0013443194,0.0013075679,0.001273081,0.0017577008,0.00044519402],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000099214936,0.00005322447,0.000654147,0.00005212015,0.000027460703,0.00005596499,0.00012287777,0.9440146,0.0008897332,0.021864386,0.0005385767,0.031627785],"study_design_scores_gemma":[0.0000061457363,0.0000065767013,0.000026852254,0.0000031373997,0.0000018286543,0.0000045778893,0.0000033827946,0.99288917,0.00018529505,0.0066901552,0.000180499,0.000002400532],"about_ca_topic_score_codex":0.00784755,"about_ca_topic_score_gemma":0.008603603,"teacher_disagreement_score":0.00784755,"about_ca_system_score_codex":0.0013494708,"about_ca_system_score_gemma":0.0020303223,"threshold_uncertainty_score":0.015603721},"labels":[],"label_agreement":null},{"id":"W2999709951","doi":"10.1109/ictai52525.2021.00026","title":"Query-Based Summarization using Reinforcement Learning and Transformer Model","year":2021,"lang":"en","type":"article","venue":"2021 IEEE 33rd International Conference on Tools with Artificial Intelligence (ICTAI)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Reinforcement learning; Computer science; Multi-document summarization; Artificial intelligence; Transformer; Sentence; Cluster analysis; Field (mathematics); Ranking (information retrieval); Unsupervised learning; Machine learning; Natural language processing; Information retrieval; Data mining; Mathematics","score_opus":0.13992104009831743,"score_gpt":0.32597460905169984,"score_spread":0.1860535689533824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2999709951","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020910423,0.0005104113,0.97569454,0.00028828945,0.00004038265,0.0001194787,0.00011897158,0.0011512602,0.0011662514],"genre_scores_gemma":[0.74177563,0.00069645257,0.24984664,0.0002930144,0.00018298884,0.00033481172,0.00093455904,0.0002540946,0.005681755],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990753,0.00032781163,0.00008040482,0.00023455764,0.00019968576,0.00008217556],"domain_scores_gemma":[0.99733317,0.0015034898,0.0003185834,0.00022928455,0.0005049037,0.00011053101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017415085,0.0008063675,0.0013441689,0.0008972821,0.0003838344,0.00090096885,0.0013350054,0.00086365064,0.0019375224],"category_scores_gemma":[0.005059141,0.00030868044,0.000865061,0.0010320398,0.00049448165,0.0022448804,0.00072045275,0.0010662936,0.00060307933],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003645923,0.0002675919,0.0010031252,0.0003375148,0.00014373723,0.00018681157,0.0003672474,0.7094032,0.015940981,0.01967025,0.0046938346,0.24762107],"study_design_scores_gemma":[0.000015759062,0.00006639105,0.000089201596,0.0000034643167,0.00001648659,0.000018215276,0.000011732906,0.99424225,0.0013084954,0.0037876328,0.00043367132,0.000006571717],"about_ca_topic_score_codex":0.0030975905,"about_ca_topic_score_gemma":0.0029606814,"teacher_disagreement_score":0.0030975905,"about_ca_system_score_codex":0.00092637294,"about_ca_system_score_gemma":0.000818125,"threshold_uncertainty_score":0.00921011},"labels":[],"label_agreement":null},{"id":"W3000138314","doi":"10.1145/3297662.3365789","title":"Toward A Real-Time Social Recommendation System","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Latent Dirichlet allocation; Conversation; Computer science; Recommender system; Topic model; Coherence (philosophical gambling strategy); World Wide Web; Information retrieval; Human–computer interaction","score_opus":0.030859961274444712,"score_gpt":0.2500916988988041,"score_spread":0.2192317376243594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3000138314","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04599789,0.0003882086,0.87950844,0.0007582616,0.00010562144,0.00048219468,0.00085627363,0.06893113,0.0029719563],"genre_scores_gemma":[0.24791849,0.0002314641,0.7405211,0.00031295195,0.00012975803,0.000393346,0.0020814256,0.0006706789,0.0077407905],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997926,0.0006637398,0.00019082408,0.0005335598,0.0005675264,0.00011824719],"domain_scores_gemma":[0.9947207,0.0016995686,0.0003727745,0.0012380956,0.0015935454,0.00037524645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004530298,0.0010144424,0.0012768528,0.0021628395,0.0011072467,0.0020934907,0.0024837223,0.0023118388,0.0033387535],"category_scores_gemma":[0.0074954336,0.0007710882,0.000940614,0.001556227,0.0004084724,0.0032116491,0.0013245883,0.0014744442,0.0041396827],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024144696,0.0012253178,0.015099393,0.00059132825,0.00055848784,0.0010470154,0.0015311259,0.04822961,0.071331955,0.009376484,0.047617123,0.8009776],"study_design_scores_gemma":[0.00011829243,0.00026382192,0.0021383315,0.000031743628,0.00013796458,0.00030682868,0.00026795658,0.9584952,0.016218038,0.0046722675,0.017253675,0.00009587692],"about_ca_topic_score_codex":0.012317802,"about_ca_topic_score_gemma":0.012902702,"teacher_disagreement_score":0.012317802,"about_ca_system_score_codex":0.0007323948,"about_ca_system_score_gemma":0.0009564329,"threshold_uncertainty_score":0.024492204},"labels":[],"label_agreement":null},{"id":"W3002533498","doi":"10.22148/001c.11775","title":"Annotating Narrative Levels: Review of Guideline No. 7","year":2020,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narratology; Narrative; Guideline; Grammar; Interpretation (philosophy); Hierarchy; Focus (optics); Linguistics; Plot (graphics); Computer science; Scope (computer science); Sociology; Philosophy; Political science; Mathematics; Programming language","score_opus":0.10099942756746239,"score_gpt":0.3415630571593329,"score_spread":0.2405636295918705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3002533498","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031638478,0.6811354,0.08070232,0.13382772,0.035570454,0.003793259,0.007151289,0.001807017,0.05284873],"genre_scores_gemma":[0.01888139,0.61020046,0.2058755,0.08909656,0.0071760667,0.010484939,0.021583952,0.0015942943,0.035106864],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98255527,0.005136233,0.005536424,0.001603623,0.004554731,0.0006138338],"domain_scores_gemma":[0.91551244,0.034844037,0.0058221426,0.004597663,0.03757048,0.0016531962],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.027127584,0.001083226,0.0025881825,0.010695815,0.0017123871,0.0045349193,0.0099614,0.0072881905,0.007290058],"category_scores_gemma":[0.103049524,0.0013295612,0.002418775,0.0066224756,0.0040661558,0.007020021,0.004804173,0.005187756,0.008010111],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008105247,0.000079394165,0.00053233,0.053795125,0.00012215669,0.00038363985,0.0020891682,0.0005289004,0.0016453845,0.037086967,0.45456964,0.4490862],"study_design_scores_gemma":[0.000018651597,0.00003070955,0.0006679945,0.05905139,0.00011800985,0.0002630913,0.00040290676,0.00013137145,0.00052625604,0.006826656,0.9319293,0.00003357982],"about_ca_topic_score_codex":0.021348042,"about_ca_topic_score_gemma":0.035197355,"teacher_disagreement_score":0.97287244,"about_ca_system_score_codex":0.0052152798,"about_ca_system_score_gemma":0.024763368,"threshold_uncertainty_score":0.14346606},"labels":[],"label_agreement":null},{"id":"W3004630545","doi":"10.48550/arxiv.1911.02085","title":"Path-Based Contextualization of Knowledge Graphs for Textual Entailment","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Contextualization; Textual entailment; Path (computing); Logical consequence; Computer science; Knowledge graph; Natural language processing; Artificial intelligence; Linguistics; Philosophy; Programming language; Interpretation (philosophy)","score_opus":0.08607791740703345,"score_gpt":0.21072855744242092,"score_spread":0.12465064003538746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3004630545","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009297658,0.0004675122,0.9848947,0.00026518962,0.000024482177,0.00018439742,0.0011498492,0.0025977248,0.0011184043],"genre_scores_gemma":[0.17100868,0.0006454557,0.82084996,0.00016208758,0.00009756844,0.00034435283,0.0055837645,0.0004528447,0.0008552023],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981098,0.00063843717,0.00012351977,0.000705388,0.0003370349,0.00008578606],"domain_scores_gemma":[0.99545455,0.0027741394,0.00037688162,0.0008869579,0.00040513722,0.000102322985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015320668,0.0009954115,0.0008201927,0.004224224,0.00092054997,0.0013009259,0.001264004,0.0011324347,0.0057845423],"category_scores_gemma":[0.011506966,0.00059191225,0.0018756909,0.0036769488,0.00088909554,0.004326561,0.0022008533,0.0014352922,0.0015111738],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005931897,0.00030497654,0.0048528262,0.001477902,0.0003463782,0.0006519016,0.0014733355,0.17798315,0.020341797,0.13710368,0.01775421,0.6371166],"study_design_scores_gemma":[0.000056373254,0.00011595846,0.0013471207,0.00013168142,0.00018539338,0.00032686992,0.00034524608,0.7411674,0.01233219,0.22110851,0.022835752,0.0000475123],"about_ca_topic_score_codex":0.0044933967,"about_ca_topic_score_gemma":0.007847089,"teacher_disagreement_score":0.0057845423,"about_ca_system_score_codex":0.0009724123,"about_ca_system_score_gemma":0.0013624077,"threshold_uncertainty_score":0.019351244},"labels":[],"label_agreement":null},{"id":"W3005700414","doi":"10.1016/j.ins.2020.02.040","title":"Plausibility-promoting generative adversarial network for abstractive text summarization with multi-task constraint","year":2020,"lang":"en","type":"article","venue":"Information Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Shenzhen Institutes of Advanced Technology Innovation Program for Excellent Young Researchers; Natural Science Foundation of Guangdong Province; National Natural Science Foundation of China","keywords":"Computer science; Automatic summarization; Artificial intelligence; Discriminative model; Natural language processing; Generative grammar; Discriminator; Generative model; Classifier (UML); Inference; Adversarial system; Task (project management); Coreference; Machine learning; Resolution (logic)","score_opus":0.051649798074390754,"score_gpt":0.2774656492488205,"score_spread":0.22581585117442976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3005700414","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009541461,0.0007833101,0.98637414,0.00053788413,0.00010337425,0.00007068914,0.0002807491,0.0010134577,0.0012949299],"genre_scores_gemma":[0.6029726,0.0011656712,0.37127814,0.001100701,0.0007042028,0.0006133041,0.0032104144,0.0007139566,0.01824094],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991584,0.00034234993,0.000048513975,0.00024783608,0.00012080986,0.00008212818],"domain_scores_gemma":[0.9970912,0.002244147,0.00016615227,0.00017601691,0.00023765603,0.000084842126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017854687,0.0014712212,0.0015089357,0.0008871463,0.00062467804,0.0010452201,0.0021715786,0.002285932,0.00365836],"category_scores_gemma":[0.0056564854,0.00075386564,0.0011576876,0.0011063489,0.000835804,0.0018803485,0.0019169799,0.0027774468,0.0014687253],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050475966,0.00014355786,0.00047584626,0.00040596383,0.00019012633,0.00027485634,0.0002805092,0.73319817,0.010583708,0.019104552,0.010372545,0.22446536],"study_design_scores_gemma":[0.00001068446,0.000023121303,0.000051973246,0.000009978063,0.000018147202,0.000017430619,0.0000076020597,0.9916324,0.0008561866,0.006827339,0.0005384813,0.0000066604725],"about_ca_topic_score_codex":0.0032001128,"about_ca_topic_score_gemma":0.0052708304,"teacher_disagreement_score":0.00365836,"about_ca_system_score_codex":0.00075905764,"about_ca_system_score_gemma":0.0010497276,"threshold_uncertainty_score":0.012238443},"labels":[],"label_agreement":null},{"id":"W3006529423","doi":"10.48550/arxiv.2002.07397","title":"Improving Multi-Turn Response Selection Models with Complementary Last-Utterance Selection by Instance Weighting","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Utterance; Task (project management); Weighting; Context (archaeology); Set (abstract data type); Selection (genetic algorithm); Artificial intelligence; Noise (video); Machine learning; Conversation; Training set; Focus (optics); Natural language processing; Data mining","score_opus":0.06811481835328083,"score_gpt":0.19001280178015792,"score_spread":0.12189798342687709,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3006529423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.109917715,0.002370925,0.87735623,0.0010449729,0.00021809406,0.0003084548,0.0005343283,0.0053878697,0.0028614022],"genre_scores_gemma":[0.82324386,0.00049748336,0.16508125,0.00079307915,0.0003770233,0.0004961512,0.0019141785,0.00057035766,0.0070265736],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998204,0.0009993168,0.000070028014,0.0004195655,0.00015043176,0.00015674652],"domain_scores_gemma":[0.99560046,0.0031961491,0.00015987552,0.0003820805,0.00047091726,0.00019050036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045348913,0.0020308867,0.0020428086,0.0012408752,0.0005775611,0.0014846595,0.0027516638,0.0019158398,0.0025752736],"category_scores_gemma":[0.008240664,0.00069211086,0.0014686226,0.00091471523,0.0006257563,0.002372239,0.001413944,0.0029397113,0.0020683617],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021754014,0.00095839833,0.0068687033,0.0003791636,0.00062984164,0.00029879165,0.00088127586,0.43529958,0.018656751,0.0069160983,0.015627565,0.51130843],"study_design_scores_gemma":[0.00003289394,0.000054244367,0.00029570487,0.0000082019615,0.00004318206,0.000025410975,0.000038442242,0.9948744,0.001602078,0.0024231141,0.0005874155,0.000014982641],"about_ca_topic_score_codex":0.0058178133,"about_ca_topic_score_gemma":0.0063795946,"teacher_disagreement_score":0.0058178133,"about_ca_system_score_codex":0.00085864775,"about_ca_system_score_gemma":0.0010161864,"threshold_uncertainty_score":0.023983061},"labels":[],"label_agreement":null},{"id":"W3007129127","doi":"10.18653/v1/2020.emnlp-main.713","title":"Unsupervised Question Decomposition for Question Answering","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research","funders":"Samsung Advanced Institute of Technology; Samsung; Open Philanthropy Project; Nvidia; National Science Foundation","keywords":"Question answering; Computer science; Leverage (statistics); Open domain; Artificial intelligence; Heuristic; Baseline (sea); Information retrieval; Machine learning; Natural language processing","score_opus":0.041922746402126954,"score_gpt":0.32177406531628955,"score_spread":0.2798513189141626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3007129127","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01729582,0.0005520508,0.97051316,0.00059295556,0.00008421066,0.00023775653,0.0007711309,0.007443211,0.002509791],"genre_scores_gemma":[0.27601135,0.00035327504,0.7103137,0.00057739025,0.00019881132,0.00051345926,0.0067647793,0.0006522572,0.004614943],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99676967,0.0017197011,0.0001609391,0.0008672642,0.00033182174,0.00015061839],"domain_scores_gemma":[0.9916398,0.004928872,0.0003793347,0.0018128977,0.00092797825,0.00031112888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003150446,0.0014072264,0.00082172715,0.0026379267,0.0008122284,0.001276668,0.0017621104,0.001724827,0.0061678137],"category_scores_gemma":[0.011803555,0.00054865936,0.0018608678,0.0017503409,0.001250391,0.0045615993,0.0027975156,0.0035147853,0.0036821733],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003562268,0.00073350733,0.0053788256,0.0009858896,0.00018730684,0.00012714323,0.0014503293,0.07056061,0.04440766,0.058736265,0.034272823,0.7828035],"study_design_scores_gemma":[0.000044763397,0.00015808492,0.0014214211,0.000055813707,0.000043151296,0.00017632116,0.00023104991,0.8413366,0.014424475,0.12628223,0.01579398,0.000032181102],"about_ca_topic_score_codex":0.0021173973,"about_ca_topic_score_gemma":0.0033358685,"teacher_disagreement_score":0.0061678137,"about_ca_system_score_codex":0.00123354,"about_ca_system_score_gemma":0.0013288953,"threshold_uncertainty_score":0.02063334},"labels":[],"label_agreement":null},{"id":"W3009019702","doi":"10.1007/978-3-030-42835-8_10","title":"Cross-Level Matching Model for Information Retrieval","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Matching (statistics); Information retrieval; Artificial intelligence; Statistics; Mathematics","score_opus":0.045019070994389516,"score_gpt":0.27833373218575,"score_spread":0.23331466119136052,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3009019702","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0087166345,0.0028331496,0.983343,0.00039880513,0.00016016595,0.00004805295,0.0005382686,0.0013994252,0.0025624144],"genre_scores_gemma":[0.4443248,0.004640436,0.48497227,0.0007304963,0.0005921678,0.00046389332,0.005409931,0.000928888,0.05793713],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99868387,0.00052984204,0.000094394534,0.00031280814,0.0002407614,0.00013837055],"domain_scores_gemma":[0.9983052,0.00095348683,0.00009733628,0.00033158227,0.00026341894,0.000049018257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028118466,0.00072554586,0.0019454114,0.001338871,0.00049290113,0.0017210246,0.0034201515,0.0023552007,0.010814857],"category_scores_gemma":[0.0051935525,0.0005099664,0.0016261264,0.0029274945,0.0005437554,0.0038147876,0.0014413155,0.0019229458,0.006226409],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063385704,0.0004317619,0.0019217639,0.0006456257,0.0004710125,0.0002227315,0.00018893804,0.25385615,0.007092749,0.110713206,0.026299797,0.5975224],"study_design_scores_gemma":[0.000013188397,0.000048045058,0.00036008633,0.0000114129325,0.00006161341,0.00006664285,0.000011375389,0.96610725,0.0007731459,0.029594032,0.0029388885,0.000014308567],"about_ca_topic_score_codex":0.0058266735,"about_ca_topic_score_gemma":0.004092312,"teacher_disagreement_score":0.010814857,"about_ca_system_score_codex":0.0010569849,"about_ca_system_score_gemma":0.0009344946,"threshold_uncertainty_score":0.036179245},"labels":[],"label_agreement":null},{"id":"W3009469321","doi":"10.1145/3377325.3377513","title":"NJM-Vis","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of the Fraser Valley; University of British Columbia","funders":"","keywords":"Interpretability; Computer science; Artificial intelligence; Joint (building); Interface (matter); Task (project management); Domain (mathematical analysis); Interpretation (philosophy); Thriving; Natural language processing; Machine learning; Engineering; Psychology","score_opus":0.04730662541987404,"score_gpt":0.23831228859865597,"score_spread":0.19100566317878193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3009469321","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011670398,0.001052872,0.43014434,0.0016092274,0.0010402773,0.00034798466,0.056658205,0.41738757,0.08008909],"genre_scores_gemma":[0.12716913,0.0010968139,0.6142957,0.0021348991,0.0003685651,0.0013562521,0.097260736,0.074447505,0.08187038],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99928975,0.00015401191,0.00004991416,0.00023709632,0.00021958839,0.000049634287],"domain_scores_gemma":[0.9986607,0.00064448704,0.000058467962,0.0003514497,0.00021709235,0.00006772048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016555208,0.0013624863,0.0008442686,0.0012371235,0.00062146125,0.002814676,0.0023398786,0.0021331962,0.12168139],"category_scores_gemma":[0.008101145,0.0006062392,0.0012754683,0.00084418256,0.000462994,0.0036158846,0.0028179395,0.0023729547,0.04840408],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010662213,0.00033541297,0.0029864528,0.0010265012,0.0002791927,0.00027470608,0.0004089406,0.01896988,0.005840453,0.051563792,0.5887092,0.32853922],"study_design_scores_gemma":[0.00045177256,0.00015768621,0.0021541347,0.00025732574,0.00009945999,0.0003724635,0.00016469772,0.34141645,0.013727858,0.11899014,0.5220746,0.00013350444],"about_ca_topic_score_codex":0.0046234555,"about_ca_topic_score_gemma":0.01416187,"teacher_disagreement_score":0.12168139,"about_ca_system_score_codex":0.0009243865,"about_ca_system_score_gemma":0.0011085374,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3010385440","doi":"10.48550/arxiv.2003.03645","title":"Generating Emotionally Aligned Responses in Dialogues using Affect Control Theory","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Affect (linguistics); Conversation; Dilemma; Property (philosophy); Computer science; Control (management); Probabilistic logic; Human–computer interaction; Psychology; Cognitive psychology; Cognitive science; Artificial intelligence; Communication; Epistemology","score_opus":0.128191363884565,"score_gpt":0.21465751618846154,"score_spread":0.08646615230389654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3010385440","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09655141,0.00018902648,0.8972301,0.00038119307,0.00006590425,0.00009102647,0.00004530414,0.0003856324,0.0050603314],"genre_scores_gemma":[0.9024492,0.0001341232,0.094141126,0.00012878614,0.000036818405,0.00015916879,0.0000659595,0.00008680895,0.0027980905],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995202,0.00028555526,0.000013653787,0.00009562678,0.000057830017,0.000027144659],"domain_scores_gemma":[0.9989963,0.0007482478,0.00006908661,0.00006311053,0.00007789002,0.00004540469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009878178,0.00036596696,0.00024793608,0.00020937301,0.00024117231,0.0006978229,0.00046373252,0.0005457677,0.0018394389],"category_scores_gemma":[0.0043794475,0.00023534306,0.00046204805,0.00012293209,0.0006515647,0.00086608913,0.00085417915,0.00056541467,0.00028736706],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030294538,0.00025956787,0.003702325,0.00030668307,0.00016140238,0.00046967412,0.0029024987,0.63861495,0.07949767,0.09779922,0.0027583162,0.17322475],"study_design_scores_gemma":[0.000012571954,0.00004258874,0.00036995902,0.0000073375427,0.000012718707,0.00003052534,0.00006422391,0.96965617,0.003171807,0.025843559,0.00077980955,0.000008764385],"about_ca_topic_score_codex":0.00046898768,"about_ca_topic_score_gemma":0.00049883133,"teacher_disagreement_score":0.0018394389,"about_ca_system_score_codex":0.00034608465,"about_ca_system_score_gemma":0.00020788409,"threshold_uncertainty_score":0.006153524},"labels":[],"label_agreement":null},{"id":"W3011218358","doi":"10.2196/17652","title":"Temporal Expression Classification and Normalization From Chinese Narrative Clinical Texts: Pattern Learning Approach","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Natural Science Foundation of Guangdong Province; National Natural Science Foundation of China","keywords":"Normalization (sociology); Computer science; Artificial intelligence; Natural language processing; Feature extraction; Narrative; Expression (computer science); Pattern recognition (psychology); Machine learning; Linguistics","score_opus":0.04919639378259552,"score_gpt":0.329713661150135,"score_spread":0.28051726736753946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3011218358","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14742613,0.001793793,0.83081,0.0016405004,0.00030489187,0.0017749578,0.005403384,0.0056073116,0.005239093],"genre_scores_gemma":[0.3225312,0.0009965675,0.6610027,0.00031260244,0.00017742412,0.0012925026,0.010172436,0.0001711403,0.0033434278],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948521,0.0012581535,0.0009928193,0.0018375535,0.0008635415,0.000195862],"domain_scores_gemma":[0.98983103,0.0052448637,0.0012859523,0.0007759056,0.0027001372,0.00016214386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041307537,0.0011518878,0.00080760475,0.004192026,0.00063009595,0.0015704549,0.0013668167,0.0008022365,0.0015252429],"category_scores_gemma":[0.01637579,0.00027795357,0.0011721507,0.0035204145,0.00080334337,0.0021936023,0.0009578336,0.0012218329,0.00093309593],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004455491,0.00027273406,0.018190634,0.0007751311,0.00014115198,0.000693226,0.0012613368,0.00908036,0.020630585,0.0026963952,0.0111558335,0.934657],"study_design_scores_gemma":[0.00016006456,0.00040461918,0.039417207,0.00036867987,0.00043462426,0.001679954,0.0020427236,0.84217095,0.064599395,0.015914226,0.03265064,0.00015691624],"about_ca_topic_score_codex":0.0056945155,"about_ca_topic_score_gemma":0.0057785227,"teacher_disagreement_score":0.0056945155,"about_ca_system_score_codex":0.0014052571,"about_ca_system_score_gemma":0.0025632642,"threshold_uncertainty_score":0.021845758},"labels":[],"label_agreement":null},{"id":"W3011993190","doi":"10.1109/iiswc47752.2019.9041972","title":"Deep Learning Language Modeling Workloads: Where Time Goes on Graphics Processors","year":2019,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Softmax function; Computer science; Graphics; Language model; Vocabulary; Transformer; Deep learning; Layer (electronics); Computation; Parallel computing; Artificial intelligence; Programming language; Operating system","score_opus":0.009532136131513099,"score_gpt":0.2252336733998364,"score_spread":0.2157015372683233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3011993190","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8937341,0.0024627894,0.06527808,0.0052740066,0.00083164783,0.000111527144,0.0034604704,0.0143802585,0.014467186],"genre_scores_gemma":[0.96237797,0.00080839905,0.024997523,0.0008299298,0.00009394874,0.0000890706,0.005842003,0.0014236646,0.0035375112],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984894,0.00023338594,0.000060177532,0.00044517135,0.00044921518,0.0003225604],"domain_scores_gemma":[0.99709105,0.0014359128,0.00012070934,0.00054167915,0.00056702347,0.00024358043],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013321844,0.0013142732,0.0005815338,0.000639213,0.0004755294,0.0017196679,0.0012962469,0.00071611535,0.0044689104],"category_scores_gemma":[0.0101021575,0.0005209208,0.000528884,0.0015523983,0.0005518042,0.0048104716,0.00075099245,0.0025795505,0.0021993124],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034401363,0.00085005973,0.043434683,0.0007626844,0.00028568823,0.0013386805,0.0013569079,0.28982845,0.095927015,0.019483995,0.107453644,0.4358381],"study_design_scores_gemma":[0.00012440252,0.0003615536,0.011342923,0.00007944498,0.000099706594,0.00019874539,0.00083460746,0.91035855,0.040324163,0.012403062,0.02380694,0.00006595535],"about_ca_topic_score_codex":0.012199786,"about_ca_topic_score_gemma":0.016531618,"teacher_disagreement_score":0.012199786,"about_ca_system_score_codex":0.0011740878,"about_ca_system_score_gemma":0.0017627362,"threshold_uncertainty_score":0.02425754},"labels":[],"label_agreement":null},{"id":"W3013135062","doi":"","title":"Overview of the TREC 2018 Real-Time Summarization Track.","year":2018,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Track (disk drive); Multi-document summarization; Information retrieval; World Wide Web; Operating system","score_opus":0.06766606094889346,"score_gpt":0.2918718006947948,"score_spread":0.22420573974590133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3013135062","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021928681,0.10356001,0.2530893,0.016071685,0.015439098,0.0070267185,0.3362373,0.1493195,0.097327694],"genre_scores_gemma":[0.02725723,0.01996502,0.19007891,0.002338879,0.0036069793,0.0032519875,0.66820234,0.0064044106,0.07889424],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99648947,0.000830996,0.0004042947,0.0006320597,0.0012544177,0.00038872374],"domain_scores_gemma":[0.9880441,0.0014883612,0.0006590751,0.0012436185,0.0075012664,0.0010635437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008136649,0.002860575,0.0021968433,0.010019727,0.0020991056,0.004659898,0.0035702393,0.0018087599,0.031926934],"category_scores_gemma":[0.008990457,0.0010983705,0.0015798842,0.008820096,0.0004307086,0.005370573,0.0018798698,0.0026328657,0.04008277],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003056645,0.00024784292,0.0007240107,0.0018233804,0.00015522665,0.000088128814,0.00012412378,0.0017661419,0.013844005,0.00093072816,0.78211516,0.19787565],"study_design_scores_gemma":[0.00024625252,0.0008998708,0.008732851,0.0006598209,0.00043328851,0.0003638337,0.0002658882,0.01651647,0.029307734,0.0032058463,0.93910015,0.0002680014],"about_ca_topic_score_codex":0.0356416,"about_ca_topic_score_gemma":0.06005153,"teacher_disagreement_score":0.0356416,"about_ca_system_score_codex":0.002617656,"about_ca_system_score_gemma":0.0066367043,"threshold_uncertainty_score":0.10680622},"labels":[],"label_agreement":null},{"id":"W3013805123","doi":"","title":"Query and Answer Expansion from Conversation History.","year":2019,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Conversation; Computer science; Query expansion; Information retrieval; Linguistics; Philosophy","score_opus":0.029768301305714517,"score_gpt":0.22192630470486654,"score_spread":0.19215800339915204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3013805123","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14684497,0.0142875,0.7514319,0.0047807954,0.0014245132,0.00197897,0.027303755,0.019752037,0.032195628],"genre_scores_gemma":[0.69078827,0.0022752918,0.23974381,0.0004901854,0.0013838747,0.0010013246,0.043292414,0.000900784,0.02012403],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99749446,0.0011488323,0.00014317199,0.0005213657,0.00047030818,0.00022189085],"domain_scores_gemma":[0.9942871,0.0039800145,0.00017224786,0.0004300662,0.0009028879,0.00022781343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025920924,0.001085039,0.0010572593,0.0043377457,0.00084223936,0.0015518123,0.0012575621,0.0011460982,0.010880921],"category_scores_gemma":[0.012922151,0.00048485532,0.0008092202,0.0023020555,0.0003488438,0.0036010854,0.0017744149,0.0013714476,0.0060400283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024245817,0.0005346725,0.006180443,0.0009994353,0.00027468384,0.00035764952,0.0012739568,0.016933454,0.030787759,0.008126564,0.094253264,0.8378536],"study_design_scores_gemma":[0.00018163219,0.00035645705,0.008406175,0.00019031059,0.00036355222,0.00035919508,0.0008223002,0.8933335,0.020641578,0.02211376,0.05314454,0.0000869801],"about_ca_topic_score_codex":0.008273984,"about_ca_topic_score_gemma":0.011961036,"teacher_disagreement_score":0.010880921,"about_ca_system_score_codex":0.00081057346,"about_ca_system_score_gemma":0.0014943179,"threshold_uncertainty_score":0.03640032},"labels":[],"label_agreement":null},{"id":"W3014405193","doi":"10.1007/978-3-030-47358-7_50","title":"Investigating Citation Linkage as a Sentence Similarity Measurement Task Using Deep Learning","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Sentence; Citation; Task (project management); Similarity (geometry); Linkage (software); Information retrieval; Natural language processing; Artificial intelligence; Semantic similarity; Domain (mathematical analysis); World Wide Web; Image (mathematics)","score_opus":0.06807879427886686,"score_gpt":0.26585730152376397,"score_spread":0.1977785072448971,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014405193","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80276275,0.004004455,0.17855455,0.0021647746,0.0004936401,0.00012379761,0.0021394868,0.0025850367,0.007171583],"genre_scores_gemma":[0.9525535,0.00044046395,0.039887127,0.00019481295,0.00042744924,0.00007932866,0.0029110448,0.00019485867,0.0033114348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975867,0.0009852598,0.00019096526,0.00062242907,0.0004477892,0.00016687672],"domain_scores_gemma":[0.97646934,0.018148694,0.0015501776,0.0011310908,0.001926673,0.0007739339],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0032912523,0.0007215494,0.0010717696,0.0034906387,0.0010814397,0.002739429,0.0015458143,0.0024414815,0.0034668976],"category_scores_gemma":[0.02489044,0.00034520088,0.00073774764,0.0040185964,0.0004272692,0.0059852726,0.0017248361,0.0020484515,0.001487637],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003110699,0.003362001,0.09405171,0.0020117625,0.0008536254,0.0013499239,0.0018538651,0.04329427,0.072690465,0.031108715,0.04736708,0.69894594],"study_design_scores_gemma":[0.00009695498,0.00043743127,0.019251319,0.00007090727,0.00022682149,0.00042275887,0.0007275016,0.91583717,0.016338008,0.041025996,0.0055060405,0.0000591601],"about_ca_topic_score_codex":0.0023440155,"about_ca_topic_score_gemma":0.0027053258,"teacher_disagreement_score":0.9965094,"about_ca_system_score_codex":0.000718999,"about_ca_system_score_gemma":0.00083294965,"threshold_uncertainty_score":0.017405987},"labels":[],"label_agreement":null},{"id":"W3014923662","doi":"10.36227/techrxiv.12059019.v1","title":"Text Summarization and Classification of Clinical Discharge Summaries using Deep Learning","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Automatic summarization; Convolutional neural network; Computer science; Artificial intelligence; Sample (material); Natural language processing; Artificial neural network; Deep learning; Machine learning; Pattern recognition (psychology); Chemistry","score_opus":0.1463046939501471,"score_gpt":0.35890553475987436,"score_spread":0.21260084080972724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014923662","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27319095,0.0044486816,0.64261323,0.0038413117,0.0011846938,0.00063271576,0.025806908,0.04162565,0.006655928],"genre_scores_gemma":[0.6235674,0.0010667135,0.30510652,0.00044749965,0.0007916692,0.00033887022,0.058309786,0.00050856377,0.00986299],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992964,0.00019489661,0.00008518488,0.00018877746,0.00014877619,0.00008594482],"domain_scores_gemma":[0.99809355,0.00074969424,0.0002759103,0.00027147503,0.00051012536,0.00009924326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010235397,0.0010862211,0.0006504373,0.0025328852,0.00032517625,0.0010525775,0.00080282334,0.00077606854,0.0020908117],"category_scores_gemma":[0.0042609386,0.00023675011,0.0005867307,0.001444394,0.00019049124,0.001184174,0.00062528055,0.0011808892,0.0019388904],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000503838,0.00040116982,0.005048401,0.00043624634,0.00013446658,0.00020867688,0.00028760353,0.03564021,0.028547792,0.0019544435,0.038043525,0.88879377],"study_design_scores_gemma":[0.000062397536,0.00026318137,0.0070384126,0.00007214011,0.00008492549,0.00010263047,0.00019361374,0.93476164,0.034954313,0.0085843345,0.013834141,0.00004832332],"about_ca_topic_score_codex":0.003403874,"about_ca_topic_score_gemma":0.0071529006,"teacher_disagreement_score":0.003403874,"about_ca_system_score_codex":0.00068167766,"about_ca_system_score_gemma":0.00091889815,"threshold_uncertainty_score":0.006994486},"labels":[],"label_agreement":null},{"id":"W3015071427","doi":"10.48550/arxiv.2004.01940","title":"Pre-Trained and Attention-Based Neural Networks for Building Noetic Task-Oriented Dialogue Systems","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Adaptation (eye); Computer science; Task (project management); Track (disk drive); Artificial neural network; Artificial intelligence; Deep neural networks; Natural language processing; Human–computer interaction; Engineering; Psychology; Systems engineering","score_opus":0.05558461328526256,"score_gpt":0.19607203699786793,"score_spread":0.14048742371260536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3015071427","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16016546,0.0034776074,0.8146067,0.0008534012,0.0004944456,0.0003324766,0.0005031861,0.01130276,0.008263969],"genre_scores_gemma":[0.8154672,0.0006781695,0.17210925,0.00037439502,0.00022924253,0.00034745262,0.0015173841,0.00043615248,0.008840848],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99897385,0.0003805043,0.00004795615,0.000385961,0.000093323586,0.000118406584],"domain_scores_gemma":[0.99839944,0.00094487204,0.00007123427,0.00015108669,0.00032759173,0.0001057301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020838631,0.0015965155,0.0007738947,0.0008163646,0.00062756863,0.0011541641,0.0017263761,0.001348797,0.0022395262],"category_scores_gemma":[0.004911149,0.0006638406,0.00093848165,0.00047260904,0.00046055828,0.0026299786,0.001714121,0.0029378892,0.0012330461],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005709634,0.0005308506,0.0015477766,0.00032085774,0.0002481445,0.00015296227,0.00065130804,0.513113,0.020900685,0.0031809239,0.0076337876,0.45114872],"study_design_scores_gemma":[0.000010728801,0.000059079695,0.00023235122,0.000011404684,0.00002270618,0.000010946136,0.00003783788,0.9946268,0.0023883984,0.0018326122,0.00075764075,0.000009443569],"about_ca_topic_score_codex":0.010766493,"about_ca_topic_score_gemma":0.016522203,"teacher_disagreement_score":0.010766493,"about_ca_system_score_codex":0.0012347003,"about_ca_system_score_gemma":0.0009997354,"threshold_uncertainty_score":0.021407664},"labels":[],"label_agreement":null},{"id":"W3015648683","doi":"10.18653/v1/2021.eacl-main.246","title":"ProFormer: Towards On-Device LSH Projection Based Transformers","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Memory footprint; Computer science; Embedding; Computation; Projection (relational algebra); Precomputation; Mobile device; Dimension (graph theory); Transformer; Word embedding; Theoretical computer science; Artificial intelligence; Speech recognition; Algorithm; Mathematics; Programming language; Engineering","score_opus":0.04861438386620129,"score_gpt":0.28930380146337836,"score_spread":0.24068941759717707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3015648683","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0126257595,0.00032263904,0.96047014,0.00024385357,0.00016173413,0.00012934855,0.00055721647,0.021517966,0.003971305],"genre_scores_gemma":[0.41400853,0.00071006827,0.56029606,0.0007396947,0.0001721372,0.00033927304,0.0032335988,0.0019125716,0.018587995],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994655,0.00010705123,0.000035344332,0.00014040786,0.00019587662,0.000055804234],"domain_scores_gemma":[0.9992317,0.0002405482,0.000034226225,0.00026681594,0.00018090106,0.000045802502],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007571241,0.00089017925,0.00066849875,0.00054358336,0.00027636046,0.0014146179,0.0019424325,0.0006714007,0.016734485],"category_scores_gemma":[0.0032279722,0.0005640956,0.000568077,0.00075253844,0.0006234893,0.004559783,0.0022484323,0.0016291892,0.007565589],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078283984,0.0002773351,0.0007862711,0.00030079516,0.0000853825,0.0002260816,0.00020245985,0.02821043,0.054908894,0.029123904,0.030873057,0.85422266],"study_design_scores_gemma":[0.00009166636,0.00023651024,0.00032567023,0.000030477677,0.000033277083,0.00027106155,0.00009876378,0.8729817,0.058504157,0.047513887,0.019873293,0.00003951623],"about_ca_topic_score_codex":0.0021860756,"about_ca_topic_score_gemma":0.004195065,"teacher_disagreement_score":0.016734485,"about_ca_system_score_codex":0.0005710759,"about_ca_system_score_gemma":0.00095707097,"threshold_uncertainty_score":0.05598241},"labels":[],"label_agreement":null},{"id":"W3015713034","doi":"10.1007/978-3-030-45442-5_4","title":"Which BM25 Do You Mean? A Large-Scale Reproducibility Study of Scoring Variants","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Reproducibility; Scale (ratio); Information retrieval; Artificial intelligence; Statistics; Mathematics; Cartography","score_opus":0.03285745192629517,"score_gpt":0.2676645239402604,"score_spread":0.23480707201396525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3015713034","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9182201,0.012483212,0.055914737,0.0012191347,0.0010565085,0.00030335534,0.0038836687,0.0017946945,0.0051246295],"genre_scores_gemma":[0.9758304,0.0005911762,0.016866695,0.00020552054,0.00022451425,0.00014879463,0.0041267686,0.00076687854,0.0012391909],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9495412,0.030376296,0.0033884353,0.009881474,0.0061611496,0.0006514071],"domain_scores_gemma":[0.6871519,0.25075617,0.009507655,0.032053534,0.018065212,0.0024655832],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07356825,0.0013266425,0.0017396541,0.003176496,0.0017024194,0.003960364,0.0027430893,0.0018472422,0.0019933728],"category_scores_gemma":[0.25292847,0.00056040863,0.0022577452,0.0034561092,0.0022606484,0.0030767357,0.002014762,0.0022393249,0.0016816461],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0090097515,0.0007260534,0.5411126,0.0014579312,0.015400277,0.00037479552,0.0047753626,0.00928698,0.006855055,0.0032347701,0.038873672,0.3688927],"study_design_scores_gemma":[0.0013695877,0.004207273,0.80384076,0.0006396945,0.010202762,0.0031204077,0.0029254998,0.12420593,0.013582015,0.015376352,0.019807149,0.00072255055],"about_ca_topic_score_codex":0.0030101216,"about_ca_topic_score_gemma":0.00361878,"teacher_disagreement_score":0.9264318,"about_ca_system_score_codex":0.0010944584,"about_ca_system_score_gemma":0.000851398,"threshold_uncertainty_score":0.38907075},"labels":[{"model":"gemma","categories":["metaresearch"],"domain":"reproducibility","study_design":"simulation_or_modeling","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["metaresearch"],"domain":"reproducibility","study_design":"design_other","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W3015978110","doi":"10.48550/arxiv.2004.05707","title":"VGCN-BERT: Augmenting BERT with Graph Embedding for Text Classification","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Sentence; Graph; Vocabulary; Embedding; Convolutional neural network; Artificial intelligence; Natural language processing; Information retrieval; Theoretical computer science","score_opus":0.127657744654395,"score_gpt":0.21160187845791617,"score_spread":0.08394413380352117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3015978110","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05245585,0.0035658495,0.9141757,0.0010527149,0.00054020394,0.00030873468,0.0033799328,0.018236479,0.0062845135],"genre_scores_gemma":[0.576983,0.0022435063,0.3790077,0.00084400096,0.0005910897,0.00043129723,0.018994344,0.0013918407,0.01951332],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99950373,0.00013051085,0.000025057101,0.0001629644,0.000113483686,0.00006434291],"domain_scores_gemma":[0.99914014,0.00036619135,0.00008473431,0.00018684789,0.0001618927,0.000060245704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008372886,0.0018846913,0.0009922188,0.0033895706,0.0006007511,0.00085996726,0.002133711,0.0016193652,0.003279063],"category_scores_gemma":[0.0025263496,0.000471919,0.001053492,0.003422984,0.00056433707,0.0035134028,0.0011687379,0.0017918395,0.0023099335],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028870895,0.00032646288,0.0027930948,0.00035636875,0.00019226404,0.0001981195,0.00018817598,0.20010297,0.010066017,0.016337644,0.049070355,0.72007984],"study_design_scores_gemma":[0.000013220949,0.000031820415,0.0003717291,0.000020172502,0.00001904365,0.000040786006,0.000023527029,0.9784345,0.0013897069,0.015109984,0.0045333975,0.000012186657],"about_ca_topic_score_codex":0.017172702,"about_ca_topic_score_gemma":0.028256193,"teacher_disagreement_score":0.017172702,"about_ca_system_score_codex":0.001412692,"about_ca_system_score_gemma":0.00090806215,"threshold_uncertainty_score":0.034145474},"labels":[],"label_agreement":null},{"id":"W3016930808","doi":"10.18653/v1/2020.emnlp-main.350","title":"Do sequence-to-sequence VAEs learn global features of sentences?","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Memorization; Computer science; Sequence (biology); Autoencoder; Artificial intelligence; Sentence; Word (group theory); Natural language processing; Language model; Natural language; Deep learning; Autoregressive model; Topic model; Speech recognition; Linguistics; Mathematics","score_opus":0.08953561004095345,"score_gpt":0.3284753241173212,"score_spread":0.23893971407636777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3016930808","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20201044,0.0015275507,0.78686047,0.0014631386,0.00026250997,0.00011192289,0.00048741646,0.002208559,0.0050679063],"genre_scores_gemma":[0.90939325,0.0006774768,0.0824297,0.00051185983,0.00014586379,0.00014627053,0.0010004508,0.00034018973,0.005354935],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992461,0.0003792697,0.000027331987,0.00020141786,0.00006234204,0.000083654864],"domain_scores_gemma":[0.9941533,0.0036689932,0.00040620615,0.0011015082,0.00043791404,0.00023199225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026238593,0.0008671647,0.00070280436,0.0004961338,0.00030228065,0.0010402714,0.0010498436,0.0012654273,0.00352052],"category_scores_gemma":[0.016524388,0.00055011513,0.0005887872,0.000412853,0.0006810104,0.004688854,0.00097528304,0.0019693316,0.002084479],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007820005,0.00027591712,0.0141846705,0.00039746432,0.00045052182,0.00025242093,0.00066440157,0.3073356,0.026742632,0.034204498,0.010746711,0.6039632],"study_design_scores_gemma":[0.000030785694,0.00021613219,0.0026241639,0.000043525353,0.00004503222,0.00014857508,0.00010412839,0.94695324,0.0060252333,0.04125629,0.002522379,0.000030493287],"about_ca_topic_score_codex":0.0014085277,"about_ca_topic_score_gemma":0.0026422497,"teacher_disagreement_score":0.00352052,"about_ca_system_score_codex":0.0004024676,"about_ca_system_score_gemma":0.0005856915,"threshold_uncertainty_score":0.013876438},"labels":[],"label_agreement":null},{"id":"W3017344042","doi":"10.1016/j.fsisyn.2020.03.004","title":"To what extent if any has Twitter disrupted hierarchies in forensic pathology?","year":2020,"lang":"en","type":"editorial","venue":"Forensic Science International Synergy","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Horizon Health Network; Saint John Regional Hospital","funders":"","keywords":"Forensic pathology; Forensic science; Internet privacy; Pathology; Computer science; Biology; Medicine; Genetics; Autopsy","score_opus":0.023315639973430818,"score_gpt":0.28292109911199564,"score_spread":0.2596054591385648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3017344042","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000114996765,0.0075801546,0.00026678608,0.10790166,0.8823956,0.000015245385,0.00016519957,0.000042710813,0.0015175317],"genre_scores_gemma":[0.0013876122,0.0077655786,0.00020905145,0.019908736,0.96424073,0.00003191631,0.00008749795,0.00004811596,0.0063206893],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961294,0.001032502,0.000491027,0.00052214245,0.0016112673,0.00021368972],"domain_scores_gemma":[0.95457536,0.029391883,0.0016838571,0.0007743069,0.010407893,0.0031667552],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.011910705,0.0016719971,0.0019750888,0.003995501,0.0030769445,0.008350056,0.0023091696,0.014717069,0.0104188295],"category_scores_gemma":[0.04990014,0.000861416,0.0016913419,0.0020246105,0.0026198267,0.004387743,0.0015476391,0.012961787,0.005558448],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002678825,0.0000061775654,0.00006807442,0.00024835794,0.000017820774,0.000042102776,0.00002917715,0.000026039246,0.000020987552,0.00055960997,0.99026185,0.008692921],"study_design_scores_gemma":[0.000059839414,0.00003249555,0.0008849239,0.0014796971,0.00016388613,0.00018860355,0.00032174002,0.0005218448,0.0001301061,0.005253141,0.9909286,0.000035171906],"about_ca_topic_score_codex":0.0044159754,"about_ca_topic_score_gemma":0.013001551,"teacher_disagreement_score":0.996923,"about_ca_system_score_codex":0.0031947799,"about_ca_system_score_gemma":0.0047123,"threshold_uncertainty_score":0.062990606},"labels":[],"label_agreement":null},{"id":"W3017372818","doi":"10.18653/v1/2020.acl-main.56","title":"Gated Convolutional Bidirectional Attention-based Model for Off-topic Spoken Response Detection","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Recall; Task (project management); Relevance (law); Residual; Pattern recognition (psychology); Sampling (signal processing); Training set","score_opus":0.04774598965152945,"score_gpt":0.2549896021022108,"score_spread":0.20724361245068132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3017372818","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14062387,0.001757934,0.84750944,0.0005394752,0.00020592993,0.00009124727,0.0007150502,0.0039168387,0.0046401736],"genre_scores_gemma":[0.93178076,0.00059374986,0.05387451,0.0003113079,0.000088217865,0.00015419092,0.0011636727,0.00016318251,0.011870407],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997595,0.00006199446,0.000010714661,0.00007736703,0.000036526435,0.00005390397],"domain_scores_gemma":[0.99954385,0.00020516894,0.000041224648,0.00005240444,0.00012702131,0.000030342293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00061278784,0.0009822891,0.0006456268,0.0004949421,0.00021321297,0.00042917122,0.0011923811,0.00064127,0.002053939],"category_scores_gemma":[0.0014439697,0.00029152204,0.00055675633,0.00045554055,0.00031323396,0.0007119146,0.0006951599,0.0011751177,0.000976115],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011546764,0.0003948086,0.006115434,0.00024929398,0.00019941575,0.00028748368,0.00035631945,0.40265715,0.064343385,0.008565705,0.010798115,0.50487816],"study_design_scores_gemma":[0.000007752932,0.000033338416,0.0005820273,0.000006195068,0.000023927447,0.000025013076,0.000009188993,0.9941592,0.0032386235,0.001307142,0.00060021796,0.000007310122],"about_ca_topic_score_codex":0.011805636,"about_ca_topic_score_gemma":0.014776553,"teacher_disagreement_score":0.011805636,"about_ca_system_score_codex":0.0006887117,"about_ca_system_score_gemma":0.0009779955,"threshold_uncertainty_score":0.0234738},"labels":[],"label_agreement":null},{"id":"W3017779903","doi":"10.18653/v1/2020.findings-emnlp.109","title":"Quantifying the Contextualization of Word Representations with Semantic Class Probing","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Contextualization; Computer science; Natural language processing; Task (project management); Artificial intelligence; Context (archaeology); Inference; Word (group theory); Layer (electronics); Embedding; Class (philosophy); Linguistics","score_opus":0.1208758373613004,"score_gpt":0.3257294327598204,"score_spread":0.20485359539852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3017779903","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34150338,0.00234918,0.6459735,0.0010660605,0.00014730923,0.00007277448,0.00075456966,0.0049776672,0.0031555872],"genre_scores_gemma":[0.9213752,0.0006342342,0.07478374,0.00025342972,0.00008475195,0.00009232855,0.0011996889,0.00040420343,0.0011724296],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99935967,0.00016283,0.000032254102,0.00031481963,0.00005145822,0.000078962556],"domain_scores_gemma":[0.99796104,0.0011260252,0.00019375498,0.00049037294,0.0001372344,0.000091453054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009083318,0.0015347695,0.00084621366,0.0007339977,0.00040472066,0.0013383745,0.0008440872,0.0011754172,0.001876875],"category_scores_gemma":[0.0069535743,0.0006628816,0.00091944064,0.00083061395,0.0010757031,0.004636707,0.0018636829,0.0028470533,0.0008699989],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007237858,0.00031515892,0.013410125,0.0007099043,0.00047653378,0.0002509993,0.0011011441,0.36030763,0.13200755,0.017242001,0.0068194075,0.46663573],"study_design_scores_gemma":[0.000031170333,0.00012507816,0.0032499647,0.00004137819,0.00012417288,0.00006650209,0.0001351105,0.9277416,0.01941899,0.04691334,0.0021137872,0.000038895425],"about_ca_topic_score_codex":0.0028797965,"about_ca_topic_score_gemma":0.0049923453,"teacher_disagreement_score":0.0028797965,"about_ca_system_score_codex":0.00063812046,"about_ca_system_score_gemma":0.00077139464,"threshold_uncertainty_score":0.0062787533},"labels":[],"label_agreement":null},{"id":"W3018695736","doi":"10.18653/v1/2020.acl-srw.16","title":"Considering Likelihood in NLP Classification Explanations with Occlusion and Language Modeling","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Computer science; Artificial intelligence; Natural language processing; Language model; Context (archaeology); Occlusion","score_opus":0.0539578300993899,"score_gpt":0.2734041729442856,"score_spread":0.2194463428448957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3018695736","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010373831,0.00042106633,0.98685986,0.00087467785,0.000030628977,0.00007019947,0.00015653697,0.0005568586,0.0006563275],"genre_scores_gemma":[0.47626632,0.0007465573,0.51779515,0.0005199919,0.00045684818,0.0004661429,0.001367635,0.00045813218,0.0019231164],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99036026,0.005591438,0.0005018401,0.0015246707,0.0016007567,0.00042095702],"domain_scores_gemma":[0.93061185,0.060486134,0.0029943378,0.0033964335,0.0018801714,0.0006311151],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011568859,0.0012950901,0.0018963505,0.003801161,0.001600632,0.0044095563,0.0030006545,0.0037120923,0.003222651],"category_scores_gemma":[0.07256941,0.0012731684,0.0018217944,0.0033851748,0.0034474323,0.009158522,0.0048943223,0.0045602545,0.0006297196],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062658085,0.00021864606,0.023114625,0.00076158374,0.0003647112,0.0011595904,0.004824877,0.30742195,0.0027731783,0.331995,0.0095703,0.31716886],"study_design_scores_gemma":[0.000041284045,0.0000396562,0.0012478877,0.000078331104,0.000054386583,0.0002353682,0.00018586867,0.76998097,0.001064021,0.2231367,0.003884563,0.000050910294],"about_ca_topic_score_codex":0.006872055,"about_ca_topic_score_gemma":0.005735136,"teacher_disagreement_score":0.011568859,"about_ca_system_score_codex":0.0025391853,"about_ca_system_score_gemma":0.0021332735,"threshold_uncertainty_score":0.061182678},"labels":[],"label_agreement":null},{"id":"W3019416653","doi":"10.18653/v1/2021.acl-long.416","title":"StereoSet: Measuring stereotypical bias in pretrained language models","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":102,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Canadian Institute for Advanced Research","funders":"","keywords":"Stereotype (UML); Race (biology); Set (abstract data type); Scale (ratio); Language model; Computer science; Measure (data warehouse); Natural language processing; Artificial intelligence; Psychology; Cognitive psychology; Social psychology; Data mining; Geography","score_opus":0.13005698811730562,"score_gpt":0.2873309265116926,"score_spread":0.157273938394387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3019416653","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7834313,0.0026839657,0.18847288,0.00055442227,0.0005632018,0.00021514678,0.009603158,0.0076321345,0.006843806],"genre_scores_gemma":[0.9386572,0.0003383627,0.04138286,0.0002095349,0.0001513136,0.00018687196,0.016501006,0.0007979414,0.0017749491],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99789727,0.0008545154,0.000118574535,0.0006324283,0.00036142694,0.00013571954],"domain_scores_gemma":[0.9905591,0.006257132,0.00050545007,0.0015371302,0.0008223476,0.00031897388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039011047,0.00081494305,0.0007968158,0.0014622394,0.0005171322,0.0011883961,0.0011930265,0.0012164181,0.0035009773],"category_scores_gemma":[0.020433774,0.00044724956,0.0005572517,0.0013740689,0.0004011138,0.0028552036,0.0017383546,0.0016904636,0.0014542893],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0076275575,0.000942898,0.1374095,0.000971775,0.002024877,0.00029680357,0.001735034,0.078888185,0.04030051,0.008536872,0.05903788,0.6622281],"study_design_scores_gemma":[0.00045317083,0.00080553215,0.0803431,0.0001227534,0.00041211396,0.00054067874,0.00077991583,0.84721226,0.030101888,0.029961694,0.009092116,0.00017474983],"about_ca_topic_score_codex":0.0038791941,"about_ca_topic_score_gemma":0.006937301,"teacher_disagreement_score":0.0039011047,"about_ca_system_score_codex":0.0006505046,"about_ca_system_score_gemma":0.00066123094,"threshold_uncertainty_score":0.020631254},"labels":[],"label_agreement":null},{"id":"W3021034694","doi":"10.18653/v1/2020.acl-main.695","title":"The Paradigm Discovery Problem","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; New York University Abu Dhabi","keywords":"Computer science; Benchmark (surveying); Task (project management); Cluster analysis; String (physics); Heuristic; Construct (python library); Artificial intelligence; Word (group theory); Code (set theory); Natural language processing; Machine learning; Programming language","score_opus":0.03907413907118272,"score_gpt":0.25613601388132207,"score_spread":0.21706187481013933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021034694","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019226111,0.0021608586,0.9545023,0.010126161,0.0006760683,0.00028263006,0.0012917634,0.0013790725,0.010354948],"genre_scores_gemma":[0.2513743,0.0026006892,0.71658385,0.005228618,0.0018523479,0.0008375871,0.007256854,0.0010330118,0.013232649],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97932386,0.0071926606,0.0015706042,0.0072624763,0.003771611,0.0008788305],"domain_scores_gemma":[0.9344056,0.047754653,0.0027739871,0.009646507,0.0040553263,0.0013638717],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015476559,0.0017759054,0.0019165584,0.0034760463,0.0042572706,0.0076245246,0.0059652915,0.0052661626,0.012069575],"category_scores_gemma":[0.0725817,0.0015299786,0.004278889,0.0036490038,0.0050457637,0.03040476,0.009727047,0.008060665,0.0041853474],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003727143,0.00023207822,0.007540875,0.0014502364,0.00037163496,0.0008739646,0.0023353947,0.010408702,0.003085402,0.49247447,0.044404574,0.43644997],"study_design_scores_gemma":[0.00004231118,0.00005782106,0.00049433985,0.00011416286,0.000068333524,0.0014089263,0.0006956406,0.047916066,0.00214119,0.91082674,0.036179606,0.00005472279],"about_ca_topic_score_codex":0.0019032095,"about_ca_topic_score_gemma":0.0014460138,"teacher_disagreement_score":0.015476559,"about_ca_system_score_codex":0.0021774792,"about_ca_system_score_gemma":0.003673289,"threshold_uncertainty_score":0.0818488},"labels":[],"label_agreement":null},{"id":"W3021107458","doi":"10.18653/v1/2020.eval4nlp-1.5","title":"BLEU Neighbors: A Reference-less Approach to Automatic Evaluation","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"BLEU; Computer science; Machine translation; Artificial intelligence; Natural language processing; Natural language generation; Evaluation of machine translation; Language model; Bottleneck; Lexical diversity; Ground truth; Sentence; Machine learning; Natural language; Vocabulary; Linguistics","score_opus":0.20922197770817,"score_gpt":0.3298302363485546,"score_spread":0.12060825864038457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021107458","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024820708,0.0030507208,0.92831117,0.00034880106,0.0004185712,0.0006355269,0.004812135,0.026033511,0.011568927],"genre_scores_gemma":[0.3791505,0.0008229235,0.58650255,0.00043590812,0.00030752236,0.0015530479,0.01560765,0.0064041293,0.009215697],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.970064,0.017561434,0.0017209072,0.004405268,0.0058034197,0.0004449144],"domain_scores_gemma":[0.9498284,0.026112957,0.0024186182,0.010799233,0.010199547,0.00064126495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020047607,0.0024895535,0.0021028034,0.008808377,0.0012811524,0.0038094653,0.0027186624,0.0026874729,0.0050462233],"category_scores_gemma":[0.0826691,0.00073443877,0.0011659308,0.0052578216,0.0011147896,0.003989181,0.002809067,0.0025784026,0.0044132317],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010916429,0.0005022953,0.01895314,0.0011709183,0.0012926124,0.00018656511,0.0010576661,0.09228644,0.011669865,0.0195004,0.06885578,0.7834326],"study_design_scores_gemma":[0.00019750197,0.0009230273,0.011938694,0.00035780764,0.00029867873,0.00047049654,0.00038183128,0.8466069,0.033565957,0.056346215,0.04857139,0.0003415344],"about_ca_topic_score_codex":0.0048322133,"about_ca_topic_score_gemma":0.009948242,"teacher_disagreement_score":0.020047607,"about_ca_system_score_codex":0.0018622626,"about_ca_system_score_gemma":0.0019135872,"threshold_uncertainty_score":0.10602313},"labels":[],"label_agreement":null},{"id":"W3021643468","doi":"10.18653/v1/2020.acl-main.679","title":"The Sensitivity of Language Models and Humans to Winograd Schema Perturbations","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Schema (genetic algorithms); Language model; Associative property; Artificial intelligence; Task (project management); Language understanding; Cognitive psychology; Synonym (taxonomy); Natural language processing; Machine learning; Psychology; Mathematics","score_opus":0.05336245230994675,"score_gpt":0.28536201352671325,"score_spread":0.23199956121676651,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021643468","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70681864,0.015348563,0.192064,0.011934686,0.0025893073,0.00050403556,0.019447083,0.023078376,0.028215248],"genre_scores_gemma":[0.9064029,0.0012154176,0.05444836,0.0025283955,0.00024282747,0.00025612544,0.029157085,0.0013732772,0.0043757386],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9875448,0.005615902,0.00068564474,0.0043897172,0.001298709,0.0004652852],"domain_scores_gemma":[0.96333337,0.023379667,0.0013437491,0.0096700955,0.0013828956,0.00089023134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013882376,0.0026982569,0.0014644454,0.0015048898,0.0010201946,0.004422647,0.0027069242,0.0025452431,0.005327607],"category_scores_gemma":[0.061114576,0.0009791972,0.0013935952,0.0014130063,0.0020621438,0.008459159,0.0040590293,0.0068426793,0.004116045],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028699236,0.0009006233,0.0723984,0.0024084765,0.0021867016,0.0009081684,0.002211591,0.17975114,0.022432275,0.015752776,0.16349138,0.53468853],"study_design_scores_gemma":[0.00041465095,0.00055493414,0.025562476,0.0004904334,0.0003614674,0.001498703,0.0018204807,0.7735078,0.026136942,0.118112415,0.051243845,0.00029589515],"about_ca_topic_score_codex":0.007368538,"about_ca_topic_score_gemma":0.010161002,"teacher_disagreement_score":0.013882376,"about_ca_system_score_codex":0.0016312319,"about_ca_system_score_gemma":0.0014490272,"threshold_uncertainty_score":0.07341796},"labels":[],"label_agreement":null},{"id":"W3021780374","doi":"10.48550/arxiv.2004.14560","title":"RikiNet: Reading Wikipedia Pages for Natural Question Answering","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Paragraph; Question answering; Computer science; Open domain; Reading (process); Dual (grammatical number); Natural (archaeology); Natural language processing; Information retrieval; Set (abstract data type); Artificial intelligence; Natural language; Domain (mathematical analysis); Questions and answers; World Wide Web; Linguistics; Programming language; Mathematics","score_opus":0.08778468336748627,"score_gpt":0.2043180002290689,"score_spread":0.11653331686158264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021780374","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09770195,0.0075426,0.605924,0.002514508,0.0018408361,0.0022767254,0.040623512,0.2224576,0.019118298],"genre_scores_gemma":[0.278814,0.0016613551,0.5746676,0.0019223982,0.00043686054,0.0019203422,0.10861814,0.0035324593,0.028426746],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987618,0.00040622958,0.00006140067,0.00046511088,0.000216318,0.00008914808],"domain_scores_gemma":[0.9977718,0.0010719203,0.000088675195,0.00047472498,0.00043242835,0.00016032292],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019234858,0.0024231167,0.0008384538,0.0018947761,0.0005697379,0.0013398961,0.0030161995,0.0021445316,0.0063903793],"category_scores_gemma":[0.0068454337,0.0005650244,0.00127087,0.00088254415,0.0004285515,0.0043620197,0.0023067235,0.0021123225,0.007892332],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073865,0.00080279476,0.004808315,0.0016950718,0.00037928738,0.00045440908,0.00074128015,0.02465117,0.028068252,0.003597892,0.25349405,0.68056893],"study_design_scores_gemma":[0.0002635555,0.000762826,0.0056217527,0.00020773533,0.00023656974,0.0007969387,0.00045486778,0.7676423,0.044373162,0.01749434,0.16194972,0.00019616456],"about_ca_topic_score_codex":0.008113772,"about_ca_topic_score_gemma":0.015524783,"teacher_disagreement_score":0.008113772,"about_ca_system_score_codex":0.0008075463,"about_ca_system_score_gemma":0.0012110763,"threshold_uncertainty_score":0.02137798},"labels":[],"label_agreement":null},{"id":"W3021792233","doi":"10.1145/3440755","title":"Evolution of Semantic Similarity—A Survey","year":2021,"lang":"en","type":"article","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":281,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Semantic similarity; Similarity (geometry); Semantic computing; Field (mathematics); Natural language; Open research; Semantics (computer science); Strengths and weaknesses; Natural language understanding","score_opus":0.04743587048223354,"score_gpt":0.2804255950414988,"score_spread":0.23298972455926523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3021792233","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03389379,0.6764765,0.24514586,0.0060678925,0.0010695993,0.00019210976,0.001122306,0.00092982466,0.035102163],"genre_scores_gemma":[0.2969955,0.5122386,0.17333005,0.0012090058,0.003599612,0.00030920733,0.0031857889,0.0005612049,0.008570908],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9954259,0.000997416,0.0004939847,0.0011478513,0.0017655914,0.0001692958],"domain_scores_gemma":[0.989225,0.006573655,0.0006124779,0.0011910915,0.0021301324,0.00026767206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042515257,0.00093792746,0.0020442691,0.009730139,0.0008492635,0.0050131776,0.0020667363,0.0016772281,0.0034430493],"category_scores_gemma":[0.018401824,0.00080037984,0.0011410796,0.012662808,0.0017225624,0.012216265,0.0024856615,0.0017858155,0.0017661721],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008821595,0.00012661226,0.0078920815,0.0020754775,0.00014977372,0.00008897273,0.0005095289,0.0050972574,0.0009662251,0.09305164,0.0101629,0.8797913],"study_design_scores_gemma":[0.000033349887,0.0003803133,0.01927441,0.0023922883,0.00028887135,0.0026356135,0.0021518576,0.09958244,0.004810144,0.42700845,0.4412337,0.00020864166],"about_ca_topic_score_codex":0.0023740027,"about_ca_topic_score_gemma":0.0013473487,"teacher_disagreement_score":0.009730139,"about_ca_system_score_codex":0.002027616,"about_ca_system_score_gemma":0.0017644865,"threshold_uncertainty_score":0.022484481},"labels":[],"label_agreement":null},{"id":"W3022187094","doi":"10.1609/aaai.v31i1.10983","title":"A Hierarchical Latent Variable Encoder-Decoder Model for Generating Dialogues","year":2017,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":702,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; McGill University; Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Samsung; Compute Canada; Samsung Advanced Institute of Technology; Canadian Institute for Advanced Research","keywords":"Computer science; Latent variable; Generative grammar; Generative model; Latent variable model; Artificial intelligence; Variable (mathematics); Artificial neural network; Encoder; Task (project management); Machine learning; Mathematics","score_opus":0.09665800747503062,"score_gpt":0.2927783749652687,"score_spread":0.1961203674902381,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022187094","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015729701,0.00025843453,0.9806574,0.00039012873,0.000056468696,0.00006978428,0.00031398004,0.0011353593,0.0013886895],"genre_scores_gemma":[0.6180487,0.00036205794,0.3707715,0.00027795468,0.00012212775,0.00054199155,0.0011782629,0.00031030242,0.00838696],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925035,0.0003807216,0.00003457821,0.00018190766,0.00009407,0.000058361056],"domain_scores_gemma":[0.99807334,0.0014631533,0.00009773751,0.000117468815,0.00018218339,0.000066149914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016552638,0.0007402865,0.0005835815,0.00061558455,0.0003768189,0.000722288,0.0014559784,0.0012015988,0.0036627215],"category_scores_gemma":[0.005099402,0.0005496513,0.00080067205,0.00071307604,0.00061232346,0.001315971,0.0009337185,0.0019213769,0.001228312],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051253574,0.00018373683,0.001935252,0.00024585906,0.00013123843,0.0002708102,0.00072419987,0.76183146,0.0105132675,0.06516249,0.0050129234,0.15347616],"study_design_scores_gemma":[0.000012039233,0.000016575701,0.00007630081,0.000004822493,0.000009237792,0.00001988621,0.000007029505,0.9926433,0.0005092242,0.006288574,0.00040724143,0.0000057847897],"about_ca_topic_score_codex":0.005200295,"about_ca_topic_score_gemma":0.008891263,"teacher_disagreement_score":0.005200295,"about_ca_system_score_codex":0.00091889815,"about_ca_system_score_gemma":0.0011377513,"threshold_uncertainty_score":0.012253046},"labels":[],"label_agreement":null},{"id":"W3022585926","doi":"10.1007/978-3-030-47358-7_11","title":"Selection Driven Query Focused Abstractive Document Summarization","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; ENCODE; Selection (genetic algorithm); Encoder; Artificial intelligence; Representation (politics); Mechanism (biology); Sequence (biology); Natural language processing; State (computer science); Information retrieval; Algorithm","score_opus":0.016625947420913747,"score_gpt":0.23858851603286704,"score_spread":0.2219625686119533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022585926","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0338581,0.010236156,0.91451615,0.0012743847,0.0011730053,0.00084068935,0.009456263,0.019671338,0.008973888],"genre_scores_gemma":[0.17875972,0.0040881205,0.738144,0.00077497365,0.0015990754,0.0006837116,0.03409498,0.001449903,0.04040551],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987966,0.0002998206,0.00013527561,0.00022454945,0.00043097607,0.00011278152],"domain_scores_gemma":[0.9976361,0.0007894754,0.00014640667,0.0002936838,0.0010540643,0.00008025607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010771864,0.0016505576,0.0018293704,0.0029844702,0.0008109621,0.0020607351,0.0014612206,0.0011203053,0.0105804475],"category_scores_gemma":[0.0029665767,0.0004385633,0.0010399038,0.0036660174,0.00031356115,0.0017626239,0.0011717058,0.0011450038,0.008003846],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091696146,0.00021869782,0.00065452524,0.0011585102,0.00021297773,0.00031225203,0.00025484772,0.010206287,0.113200806,0.0030596734,0.09008301,0.77972144],"study_design_scores_gemma":[0.0002772536,0.0011244214,0.004953724,0.00016085338,0.0010234741,0.0011383182,0.0007299172,0.6442226,0.19103405,0.013349562,0.14180417,0.00018170566],"about_ca_topic_score_codex":0.0023108579,"about_ca_topic_score_gemma":0.004348916,"teacher_disagreement_score":0.0105804475,"about_ca_system_score_codex":0.0005087354,"about_ca_system_score_gemma":0.0010603247,"threshold_uncertainty_score":0.035395145},"labels":[],"label_agreement":null},{"id":"W3022592851","doi":"10.18653/v1/2020.acl-main.220","title":"Learning an Unreferenced Metric for Online Dialogue Evaluation","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University","funders":"McGill University; Canadian Institute for Advanced Research","keywords":"Inference; Computer science; Metric (unit); Task (project management); Artificial intelligence; Domain (mathematical analysis); Natural language processing; Quality (philosophy); Machine learning; Open domain; Question answering; Mathematics","score_opus":0.22264646163265112,"score_gpt":0.3771446736430029,"score_spread":0.1544982120103518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022592851","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04026425,0.00118183,0.9488636,0.00028617118,0.00017110533,0.00026722948,0.00077486347,0.0047982186,0.0033928151],"genre_scores_gemma":[0.65023035,0.00027179648,0.34076065,0.0003025051,0.0002581159,0.0008964114,0.0036273124,0.0010241452,0.0026286913],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9702912,0.018396266,0.0016193788,0.0046982,0.004130652,0.00086441846],"domain_scores_gemma":[0.9454311,0.033476368,0.0037139475,0.00692676,0.008868887,0.0015828905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014826458,0.0025206676,0.0018449499,0.003423371,0.00095234916,0.0024477446,0.002630862,0.0025695025,0.0027223171],"category_scores_gemma":[0.07216726,0.00055185286,0.00083520054,0.0015195648,0.0013012697,0.0046495674,0.0038255611,0.0029369271,0.0020306972],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019233171,0.0010154393,0.016082311,0.00105017,0.00049340463,0.00032426167,0.0014570993,0.15500809,0.03628891,0.015396144,0.020008052,0.7509528],"study_design_scores_gemma":[0.000059583264,0.0006395678,0.0045604315,0.00008354366,0.00006362853,0.00025261877,0.00024248022,0.9590496,0.014235034,0.01548957,0.0052173077,0.00010666347],"about_ca_topic_score_codex":0.0024218664,"about_ca_topic_score_gemma":0.0032010346,"teacher_disagreement_score":0.014826458,"about_ca_system_score_codex":0.0017439406,"about_ca_system_score_gemma":0.0016626908,"threshold_uncertainty_score":0.078410745},"labels":[],"label_agreement":null},{"id":"W3022745277","doi":"10.18653/v1/2020.acl-main.177","title":"Probing Linguistic Systematicity","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Generalization; Meaning (existential); Set (abstract data type); Perspective (graphical); Natural (archaeology); Inference; Linguistics; Natural language understanding; Computer science; Natural language; Artificial intelligence; Natural language processing; Psychology; Cognitive science; Epistemology; Philosophy; History","score_opus":0.06970837947866719,"score_gpt":0.2778822689205535,"score_spread":0.2081738894418863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022745277","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68700534,0.0012097671,0.30025864,0.0023429813,0.000045090754,0.00012840713,0.00043584354,0.0013216918,0.0072521637],"genre_scores_gemma":[0.98013824,0.00015954183,0.018859085,0.0001606057,0.0000209488,0.000076256685,0.0002177177,0.00011154887,0.00025614275],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99433255,0.003163947,0.0002774064,0.0013435167,0.00067622203,0.0002063256],"domain_scores_gemma":[0.93611324,0.040894076,0.0072247176,0.011470301,0.0035661983,0.00073148555],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008377391,0.0006223732,0.00070463895,0.0015754888,0.0007910064,0.0015827533,0.00086036767,0.001082822,0.0010727951],"category_scores_gemma":[0.0687901,0.00066478335,0.0005802457,0.0011636745,0.003956119,0.00485733,0.0032979634,0.002837458,0.00022184495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010219614,0.0003937443,0.18000671,0.0015704761,0.00091895653,0.00043923553,0.01461197,0.11593527,0.119175084,0.19804461,0.0039627487,0.36391917],"study_design_scores_gemma":[0.00006384862,0.0006650719,0.0849572,0.00018381454,0.00024822276,0.00053607754,0.0017290445,0.35092348,0.041913208,0.50976914,0.008793252,0.00021768466],"about_ca_topic_score_codex":0.0016520009,"about_ca_topic_score_gemma":0.0020322672,"teacher_disagreement_score":0.008377391,"about_ca_system_score_codex":0.0012651327,"about_ca_system_score_gemma":0.0010884979,"threshold_uncertainty_score":0.04430437},"labels":[],"label_agreement":null},{"id":"W3022773562","doi":"10.1007/978-3-030-47358-7_35","title":"Query Focused Abstractive Summarization via Incorporating Query Relevance and Transfer Learning with Transformer Models","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Automatic summarization; Transformer; Artificial intelligence; Transfer of learning; Relevance (law); Relevance feedback; Natural language processing; Information retrieval; Image retrieval; Image (mathematics)","score_opus":0.017357309336434784,"score_gpt":0.21255671522419398,"score_spread":0.1951994058877592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022773562","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011376974,0.0011281216,0.98147994,0.0003113145,0.00015720811,0.00013107822,0.00033773974,0.0036327622,0.0014448036],"genre_scores_gemma":[0.44685102,0.0016403449,0.53142065,0.0005618559,0.001008221,0.00038639293,0.00466682,0.0009224133,0.012542243],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99895006,0.0002953959,0.00010406314,0.0002805624,0.00026837044,0.00010156308],"domain_scores_gemma":[0.99818254,0.00087518996,0.00012946854,0.0002608643,0.0004865541,0.000065450404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015765663,0.0013209367,0.0017097155,0.0017648243,0.0005112206,0.001605531,0.0017541469,0.0011875745,0.0048201536],"category_scores_gemma":[0.0044834954,0.00049842,0.0012434846,0.0019698893,0.00046013462,0.0036992796,0.0015227658,0.0017686127,0.0029852334],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006738356,0.0003043232,0.0005673647,0.00040505175,0.00018642769,0.00014966006,0.00022593512,0.0652157,0.036625203,0.008873969,0.017457487,0.86931497],"study_design_scores_gemma":[0.00003889658,0.00015762643,0.00027988313,0.000013144827,0.00011624906,0.00006968756,0.00004548577,0.97702765,0.010272836,0.009005805,0.00294894,0.000023720086],"about_ca_topic_score_codex":0.0026283655,"about_ca_topic_score_gemma":0.0040955935,"teacher_disagreement_score":0.0048201536,"about_ca_system_score_codex":0.0006682863,"about_ca_system_score_gemma":0.0010478937,"threshold_uncertainty_score":0.016125083},"labels":[],"label_agreement":null},{"id":"W3022810465","doi":"10.1109/micro50266.2020.00071","title":"GOBO: Quantizing Attention-Based NLP Models for Low Latency and Energy Efficient Inference","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Quantization (signal processing); Computation; Inference; Latency (audio); Language model; Computer engineering; Parallel computing; Computer hardware; Algorithm; Artificial intelligence","score_opus":0.06401963398422118,"score_gpt":0.2784773016417433,"score_spread":0.2144576676575221,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022810465","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017822476,0.00046411657,0.9738513,0.00035684794,0.000059805174,0.0000728352,0.00034337994,0.005080035,0.0019492038],"genre_scores_gemma":[0.5428099,0.0005069075,0.44938752,0.00055167,0.000090602,0.0002816205,0.0012458822,0.0008161198,0.004309665],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99969447,0.00006678696,0.000016050406,0.000078753736,0.00010618768,0.000037735965],"domain_scores_gemma":[0.9993625,0.00035816574,0.000050743845,0.00010580757,0.00009177445,0.00003105695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007428518,0.00078179844,0.00069909164,0.00063461944,0.00036957132,0.000919938,0.0017776308,0.00089976523,0.0050108074],"category_scores_gemma":[0.003708655,0.00044201323,0.0005984882,0.00065561204,0.00066959165,0.0023846931,0.0016034058,0.0015576213,0.0011146803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049869285,0.00016630968,0.0012057537,0.00038519793,0.00013996882,0.00015464911,0.00028597715,0.54033244,0.029865317,0.03277383,0.012505031,0.38168678],"study_design_scores_gemma":[0.0000109008015,0.000013177939,0.00006962147,0.000007028066,0.0000064131536,0.000009924203,0.000009573185,0.98752594,0.002258815,0.009297215,0.0007865493,0.0000048751217],"about_ca_topic_score_codex":0.01042909,"about_ca_topic_score_gemma":0.019857325,"teacher_disagreement_score":0.01042909,"about_ca_system_score_codex":0.0010331293,"about_ca_system_score_gemma":0.00102652,"threshold_uncertainty_score":0.020736754},"labels":[],"label_agreement":null},{"id":"W3023058189","doi":"10.1007/978-3-030-47358-7_37","title":"Attending Knowledge Facts with BERT-like Models in Question-Answering: Disappointing Results and Some Explanations","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Sentence; Computer science; Question answering; Natural language processing; Set (abstract data type); Artificial intelligence; Programming language","score_opus":0.033983566764943754,"score_gpt":0.2610527462836509,"score_spread":0.22706917951870714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3023058189","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040702723,0.0065511386,0.88187635,0.03932671,0.0010152342,0.00020157432,0.0008777183,0.0034970024,0.025951628],"genre_scores_gemma":[0.66418785,0.00335389,0.3013738,0.0045647193,0.0026449752,0.00027649058,0.0020453497,0.0008584656,0.020694492],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9881555,0.007525502,0.0006011586,0.0015943799,0.0016180156,0.00050543516],"domain_scores_gemma":[0.82235605,0.15986222,0.0018274541,0.010904175,0.003250772,0.0017994086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019502832,0.0015315666,0.0025691777,0.0014173561,0.002050162,0.006722209,0.0054670707,0.008330337,0.017392313],"category_scores_gemma":[0.09941261,0.0017657703,0.0036218765,0.002473515,0.0038969321,0.032220617,0.0050356067,0.010319003,0.0038979482],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015890771,0.0010541772,0.0038714705,0.0015022355,0.00048647902,0.00041399093,0.0030868012,0.08373757,0.0027765105,0.52886415,0.06801757,0.30459985],"study_design_scores_gemma":[0.00009657205,0.0001036745,0.00049947144,0.000085720785,0.00012887792,0.00018969424,0.00031406988,0.38473898,0.0012884166,0.6037899,0.008706626,0.000058014586],"about_ca_topic_score_codex":0.0057248236,"about_ca_topic_score_gemma":0.0035541453,"teacher_disagreement_score":0.019502832,"about_ca_system_score_codex":0.0025716985,"about_ca_system_score_gemma":0.0014134626,"threshold_uncertainty_score":0.10314208},"labels":[],"label_agreement":null},{"id":"W3024228313","doi":"10.2196/18417","title":"Extraction of Information Related to Drug Safety Surveillance From Electronic Health Record Notes: Joint Modeling of Entities and Relations Using Knowledge-Aware Neural Attentive Models","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Documentation; Computer science; Medical prescription; Information extraction; Sentence; Medical record; Artificial intelligence; Medicine; Medical emergency; Information retrieval; Data mining; Natural language processing","score_opus":0.03801965279698179,"score_gpt":0.2973762146673368,"score_spread":0.25935656187035505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3024228313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21857543,0.0029570197,0.75577736,0.003592906,0.00022229992,0.0005886782,0.0073052463,0.004836487,0.006144636],"genre_scores_gemma":[0.7755918,0.0012367159,0.20918107,0.0006794682,0.00014977076,0.0004188233,0.00847041,0.00008580703,0.004186104],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99946946,0.000108485874,0.00006158373,0.00023442517,0.00007581933,0.000050243674],"domain_scores_gemma":[0.99740195,0.001853747,0.00028886276,0.00012680826,0.00027246846,0.0000562463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001149491,0.0011857095,0.00056654785,0.0020924232,0.0004517438,0.0014038493,0.0014054526,0.001430658,0.0010167263],"category_scores_gemma":[0.0046691517,0.0006047088,0.0017941983,0.0013999529,0.0004124899,0.0017855793,0.001038129,0.0019383619,0.0005521099],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055521255,0.00091422256,0.02828232,0.0005746282,0.00046373112,0.001635709,0.00072163047,0.5426437,0.008351859,0.006554927,0.010130584,0.39917144],"study_design_scores_gemma":[0.0000076374545,0.000025236513,0.0015923469,0.00002690161,0.00005361951,0.000041742634,0.000035022633,0.993082,0.0011009157,0.0030836836,0.00094128965,0.000009736625],"about_ca_topic_score_codex":0.030633507,"about_ca_topic_score_gemma":0.029424306,"teacher_disagreement_score":0.030633507,"about_ca_system_score_codex":0.0018642391,"about_ca_system_score_gemma":0.001557318,"threshold_uncertainty_score":0.060910404},"labels":[],"label_agreement":null},{"id":"W3025873536","doi":"10.48550/arxiv.2005.05864","title":"Exploiting Syntactic Structure for Better Language Modeling: A Syntactic Distance Approach","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University; Canadian Institute for Advanced Research; Université de Montréal","funders":"","keywords":"Treebank; Perplexity; Computer science; Parsing; Natural language processing; Artificial intelligence; Language model; Task (project management); Representation (politics); Syntactic structure; Ground truth; Parse tree; Syntax","score_opus":0.09857051092529527,"score_gpt":0.20306891058058102,"score_spread":0.10449839965528575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3025873536","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049363777,0.00041281633,0.946127,0.00074285467,0.000046060435,0.00004723635,0.00032521205,0.000956528,0.0019785403],"genre_scores_gemma":[0.6736137,0.0005303821,0.32016218,0.00032894066,0.00015481436,0.0001711594,0.001753416,0.0005029973,0.0027823288],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998607,0.00066701934,0.00007365638,0.00036901355,0.00020768022,0.00007559981],"domain_scores_gemma":[0.9969216,0.0018351595,0.00023924457,0.00054450025,0.00035796312,0.00010156439],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021846,0.0011617993,0.0009799093,0.0019660532,0.00070566376,0.0016915922,0.0018140321,0.0015174233,0.0017839634],"category_scores_gemma":[0.0065571764,0.00061687187,0.0014420582,0.0018830512,0.0008874138,0.0049272636,0.00264229,0.0029720892,0.0010787767],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034663262,0.0003586438,0.0068222783,0.00030622736,0.00043899906,0.00026026653,0.00085647166,0.42920265,0.02440361,0.108809836,0.006995593,0.42119876],"study_design_scores_gemma":[0.0000103410675,0.000037003738,0.00043577963,0.000011623284,0.000035115383,0.000048402504,0.000031090276,0.93611926,0.002374587,0.05973311,0.0011445117,0.00001903144],"about_ca_topic_score_codex":0.0014734364,"about_ca_topic_score_gemma":0.0026808186,"teacher_disagreement_score":0.0021846,"about_ca_system_score_codex":0.0009520874,"about_ca_system_score_gemma":0.001254619,"threshold_uncertainty_score":0.011553407},"labels":[],"label_agreement":null},{"id":"W3028435993","doi":"10.5220/0009422400580067","title":"Using BERT and XLNET for the Automatic Short Answer Grading Task","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Task (project management); Grading (engineering); Artificial intelligence; Natural language processing; Engineering","score_opus":0.1021931580915548,"score_gpt":0.294745850003279,"score_spread":0.19255269191172422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3028435993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2689721,0.0042633,0.3867345,0.0047683967,0.0039793225,0.0019597907,0.07332273,0.20572774,0.05027207],"genre_scores_gemma":[0.5557544,0.0009981625,0.23894085,0.0009489546,0.001066795,0.0014496336,0.15501575,0.0035155737,0.042309918],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835795,0.0005511103,0.00013779713,0.00045746175,0.00031050047,0.00018515032],"domain_scores_gemma":[0.99618113,0.0018537623,0.00017313333,0.00039772483,0.001043299,0.00035106423],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019995074,0.0026807166,0.0010791663,0.0037979123,0.00093803194,0.002011408,0.0012810816,0.0024731301,0.023255829],"category_scores_gemma":[0.008835179,0.0005136259,0.0009900366,0.0020787586,0.00020217165,0.0044444934,0.0016667894,0.0019917835,0.013941812],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015678639,0.0007648197,0.0070282063,0.0009763134,0.00016245182,0.00036422253,0.0003482765,0.00849463,0.024124272,0.00396666,0.19481437,0.757388],"study_design_scores_gemma":[0.00044322136,0.0009104557,0.015428255,0.00025479286,0.00025212867,0.00045521132,0.001083954,0.80878747,0.040184554,0.022616856,0.10942775,0.00015532425],"about_ca_topic_score_codex":0.010818142,"about_ca_topic_score_gemma":0.01921137,"teacher_disagreement_score":0.023255829,"about_ca_system_score_codex":0.0012463228,"about_ca_system_score_gemma":0.0019153039,"threshold_uncertainty_score":0.077798486},"labels":[],"label_agreement":null},{"id":"W3028657247","doi":"","title":"A Robust Self-Learning Method for Fully Unsupervised Cross-Lingual Mappings of Word Embeddings: Making the Method Robustly Reproducible as Well","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Computer science; Robustness (evolution); Hyperparameter; Word (group theory); Artificial intelligence; Grid; Stability (learning theory); Unsupervised learning; Machine learning; Natural language processing; Mathematics","score_opus":0.0626216750026778,"score_gpt":0.3642289984244982,"score_spread":0.3016073234218204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3028657247","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006511368,0.00012636973,0.9863354,0.00009187203,0.00015450618,0.000091200825,0.00029064136,0.0057692723,0.000629361],"genre_scores_gemma":[0.12108947,0.00015599014,0.86561507,0.00026081066,0.00016196865,0.00053245516,0.0034159156,0.002832916,0.005935428],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99237317,0.0025845193,0.00056554464,0.0025907948,0.0015803009,0.00030568725],"domain_scores_gemma":[0.98805374,0.0033216928,0.00047625927,0.004236574,0.0036161977,0.00029556258],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.005924654,0.0017932549,0.0014927393,0.0023679594,0.0013168431,0.0026028906,0.0028755958,0.0025345387,0.0043305005],"category_scores_gemma":[0.020827573,0.0010409942,0.001730147,0.0021810294,0.0012473898,0.004948051,0.005627324,0.0037650273,0.007922208],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003307281,0.00036803607,0.0026542887,0.00028674616,0.00049755763,0.00015169672,0.0004079369,0.03572384,0.03514908,0.0123011945,0.021519119,0.89060974],"study_design_scores_gemma":[0.00009746222,0.00017217836,0.0017010135,0.000049094044,0.000103826584,0.00039624935,0.00018553837,0.91431046,0.041835774,0.02888353,0.0121519305,0.00011301664],"about_ca_topic_score_codex":0.003008452,"about_ca_topic_score_gemma":0.0062197684,"teacher_disagreement_score":0.99407536,"about_ca_system_score_codex":0.000619293,"about_ca_system_score_gemma":0.0025769058,"threshold_uncertainty_score":0.03133297},"labels":[],"label_agreement":null},{"id":"W3028930880","doi":"10.48550/arxiv.2003.04642","title":"A Framework for Evaluation of Machine Reading Comprehension Gold\\n Standards","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Computer science; Correctness; Paragraph; Reading comprehension; Artificial intelligence; Natural language processing; Schema (genetic algorithms); Comprehension; Ambiguity; Popularity; Set (abstract data type); Reading (process); Machine learning; Linguistics; Psychology; World Wide Web; Programming language","score_opus":0.1712827149806367,"score_gpt":0.2700363482322421,"score_spread":0.09875363325160538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3028930880","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018086318,0.00075778674,0.968177,0.0012814649,0.000078309095,0.0006797568,0.0012643752,0.0013152953,0.008359605],"genre_scores_gemma":[0.2272702,0.00023524296,0.7657512,0.00024419118,0.00009996987,0.0024128228,0.0025148701,0.0003156,0.0011559139],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.90839005,0.060744803,0.006208775,0.008753773,0.014605414,0.0012971737],"domain_scores_gemma":[0.8530696,0.09687054,0.011215216,0.018643966,0.018801302,0.0013994027],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07408886,0.0022903124,0.0020959387,0.019327762,0.00223936,0.013298212,0.0050932947,0.0037781333,0.003883698],"category_scores_gemma":[0.19441822,0.0010646173,0.002082494,0.009527605,0.0072778002,0.010387364,0.0076274946,0.003938878,0.001395277],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042761193,0.0005609824,0.0212176,0.0011183433,0.00041120293,0.0002659762,0.006820904,0.039862275,0.009000366,0.6602158,0.010796599,0.24930245],"study_design_scores_gemma":[0.00010843584,0.00056973775,0.016763587,0.00071051496,0.00014498063,0.0003294048,0.0030356774,0.34420803,0.008801797,0.59785706,0.027208172,0.00026266286],"about_ca_topic_score_codex":0.0074170986,"about_ca_topic_score_gemma":0.005318985,"teacher_disagreement_score":0.9259111,"about_ca_system_score_codex":0.0066729286,"about_ca_system_score_gemma":0.0039540906,"threshold_uncertainty_score":0.391824},"labels":[],"label_agreement":null},{"id":"W3029422798","doi":"","title":"Multi-class Multilingual Classification of Wikipedia Articles Using Extended Named Entity Tag Set","year":2019,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Categorization; Set (abstract data type); Structuring; Natural language processing; Text categorization; Class (philosophy); Artificial intelligence; German; Information retrieval; Named entity; Linguistics","score_opus":0.13625313645858228,"score_gpt":0.22823863147610698,"score_spread":0.0919854950175247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3029422798","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61059725,0.008643776,0.32522076,0.0015624989,0.002338468,0.0006867387,0.028197106,0.008636776,0.014116628],"genre_scores_gemma":[0.8396819,0.0011840157,0.104208596,0.00021684413,0.000984824,0.00039230965,0.043201964,0.0005688253,0.009560631],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99755543,0.00059028866,0.00030406617,0.0007612456,0.00046957924,0.00031934373],"domain_scores_gemma":[0.9936998,0.003101236,0.00045204145,0.0005450082,0.0018522469,0.00034967053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028379639,0.0013822148,0.0011783381,0.010184202,0.0015086613,0.0028191656,0.0011845338,0.0014187517,0.0021418342],"category_scores_gemma":[0.0063327374,0.0003233939,0.0018699354,0.005428861,0.00040609567,0.003045803,0.0019848882,0.0014519915,0.0030803955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030176756,0.0023004096,0.0954522,0.0012772738,0.0016181684,0.002046669,0.0013979556,0.020804984,0.032628123,0.0039654872,0.050998766,0.7844922],"study_design_scores_gemma":[0.00012921295,0.0004985664,0.06469949,0.00027168883,0.001361462,0.0013892622,0.0022589068,0.8561798,0.02978007,0.012502519,0.03070146,0.00022755875],"about_ca_topic_score_codex":0.0062664924,"about_ca_topic_score_gemma":0.011449621,"teacher_disagreement_score":0.010184202,"about_ca_system_score_codex":0.00087408867,"about_ca_system_score_gemma":0.0014838779,"threshold_uncertainty_score":0.015008807},"labels":[],"label_agreement":null},{"id":"W3029927342","doi":"","title":"Contextualized Embeddings based Transformer Encoder for Sentence Similarity Modeling in Answer Selection Task","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Transformer; Encoder; Computer science; Sentence; Artificial intelligence; Language model; Natural language processing; Feature selection; Selection (genetic algorithm); Engineering; Voltage","score_opus":0.057425820451554886,"score_gpt":0.309027963334743,"score_spread":0.25160214288318816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3029927342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16390093,0.0016445526,0.80471325,0.000653683,0.00066853873,0.00041708746,0.0055663637,0.016658487,0.005777076],"genre_scores_gemma":[0.7925718,0.0005342071,0.1862885,0.00023800074,0.000245496,0.00037974893,0.013337099,0.00058804377,0.005817258],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99887747,0.000384022,0.0000871991,0.00029986975,0.00021160273,0.000139867],"domain_scores_gemma":[0.99816954,0.00070214784,0.00007810594,0.00022606006,0.0007042358,0.0001200003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013768099,0.00093798595,0.0008692113,0.0013792529,0.0004474274,0.00095226936,0.001090735,0.00096348074,0.007119871],"category_scores_gemma":[0.0043638838,0.00025596688,0.000665853,0.0010542239,0.00019542391,0.0029269685,0.0013122695,0.0014954482,0.0035404693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015715662,0.0009806359,0.0055961,0.0004979068,0.000248783,0.00028684465,0.00032552367,0.020316258,0.05460699,0.009898605,0.044353332,0.8613174],"study_design_scores_gemma":[0.00011678061,0.0004947137,0.0025023578,0.000046003195,0.00021405135,0.00032756774,0.00024661774,0.9375609,0.03966217,0.01192872,0.0068509765,0.00004909856],"about_ca_topic_score_codex":0.0029359413,"about_ca_topic_score_gemma":0.0047521926,"teacher_disagreement_score":0.007119871,"about_ca_system_score_codex":0.00045650516,"about_ca_system_score_gemma":0.0014243943,"threshold_uncertainty_score":0.023818314},"labels":[],"label_agreement":null},{"id":"W3030151190","doi":"","title":"Evaluating Approaches to Personalizing Language Models","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Perplexity; Computer science; Language model; Personalization; Artificial intelligence; Natural language processing; Adaptation (eye); Word (group theory); World Wide Web","score_opus":0.3079769636386057,"score_gpt":0.34640647007214864,"score_spread":0.038429506433542926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3030151190","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47249722,0.011102872,0.4950416,0.002418332,0.0003952516,0.0018100537,0.0012021321,0.0050395085,0.010493055],"genre_scores_gemma":[0.8015619,0.0016179744,0.19186635,0.00035119406,0.00017129803,0.0006186442,0.0016719935,0.00032032977,0.0018202517],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96263033,0.02760296,0.0018488351,0.00280948,0.004513926,0.0005944825],"domain_scores_gemma":[0.7904698,0.19104068,0.002701092,0.008089342,0.0058782957,0.0018207245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04093151,0.002124774,0.0016757558,0.004737976,0.0013210456,0.0045657884,0.002397554,0.003610964,0.0020054919],"category_scores_gemma":[0.127417,0.0009409156,0.0016472097,0.0028031,0.0014662419,0.008082716,0.0033375607,0.003214137,0.0006081397],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0063347244,0.0036556716,0.031961188,0.001922498,0.0026349765,0.00016148663,0.0029271962,0.23214753,0.0059536565,0.013220659,0.007822404,0.691258],"study_design_scores_gemma":[0.0006525703,0.0014104656,0.004714962,0.00019289408,0.0014684085,0.00011963777,0.0008012542,0.95469844,0.007566678,0.025668152,0.0025953406,0.00011107546],"about_ca_topic_score_codex":0.00896824,"about_ca_topic_score_gemma":0.012189786,"teacher_disagreement_score":0.04093151,"about_ca_system_score_codex":0.0034036443,"about_ca_system_score_gemma":0.0029886658,"threshold_uncertainty_score":0.21646911},"labels":[],"label_agreement":null},{"id":"W3030231421","doi":"10.1145/3397271.3401160","title":"Analyzing and Learning from User Interactions for Search Clarification","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Computer science; Information retrieval; Ranking (information retrieval); Search engine; Presentation (obstetrics); Representation (politics); Web search query; World Wide Web","score_opus":0.12122208160451114,"score_gpt":0.33917069960310287,"score_spread":0.21794861799859172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3030231421","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.74534345,0.0012628991,0.24810784,0.0004919456,0.000041821535,0.00041810208,0.0010166167,0.0018006201,0.0015166987],"genre_scores_gemma":[0.9564039,0.00018464371,0.04080575,0.000049083585,0.000037563623,0.00018485497,0.001622683,0.00005753785,0.0006541386],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9956656,0.002425742,0.0002831576,0.0007892046,0.000626281,0.00021005594],"domain_scores_gemma":[0.962262,0.031165674,0.0023336757,0.0018866133,0.0018576371,0.00049448013],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041012554,0.0010650052,0.0009610622,0.003029459,0.00047438242,0.0014716492,0.0007068743,0.0013016553,0.0010463338],"category_scores_gemma":[0.0394302,0.00039389145,0.0009829192,0.0016667282,0.0005501426,0.0030131312,0.0010394056,0.0014877171,0.0006300657],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002360761,0.0015529947,0.18854953,0.0013652276,0.00033525174,0.00045013736,0.009904944,0.034096956,0.05832394,0.0036510152,0.0049227215,0.6944866],"study_design_scores_gemma":[0.00008109506,0.00078047183,0.10973567,0.00009541984,0.00011580935,0.00024036979,0.0014744016,0.8637616,0.014396484,0.0058741467,0.0033309388,0.00011364415],"about_ca_topic_score_codex":0.00285991,"about_ca_topic_score_gemma":0.004471731,"teacher_disagreement_score":0.0041012554,"about_ca_system_score_codex":0.0008324514,"about_ca_system_score_gemma":0.00075554923,"threshold_uncertainty_score":0.021689832},"labels":[],"label_agreement":null},{"id":"W3030512995","doi":"","title":"Evaluating Sub-word Embeddings in Cross-lingual Models.","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Word (group theory); Natural language processing; Artificial intelligence; Lexicon; Task (project management); Space (punctuation); Vocabulary; Resource (disambiguation); Linguistics","score_opus":0.08551581455221315,"score_gpt":0.37443450547402246,"score_spread":0.28891869092180933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3030512995","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.59945846,0.015724069,0.33414683,0.0026594754,0.002871639,0.0007382498,0.012568854,0.013798783,0.018033566],"genre_scores_gemma":[0.87772155,0.0017848946,0.08768651,0.0004930767,0.00041997322,0.00040943627,0.025590025,0.0009744075,0.004920041],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9900303,0.0067518214,0.0006454022,0.0013595822,0.0008760985,0.0003367616],"domain_scores_gemma":[0.9766827,0.017642228,0.00040628115,0.0021908772,0.0025204557,0.00055745756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0122037,0.0021400321,0.0014876237,0.003390763,0.00099368,0.0033982655,0.001987271,0.002124491,0.0036045534],"category_scores_gemma":[0.034351476,0.00055992434,0.0014978296,0.0028143695,0.0005946276,0.0074437065,0.0040148743,0.0029386184,0.0033044084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031178433,0.0021573766,0.032403216,0.0013802418,0.0039509004,0.00040197448,0.0011247386,0.12464195,0.0073644007,0.0066092554,0.051671304,0.7651769],"study_design_scores_gemma":[0.0002509017,0.00079652754,0.005848217,0.00019070708,0.0009837814,0.00027100267,0.0011182139,0.9616428,0.007858173,0.013392046,0.0075486708,0.000098973396],"about_ca_topic_score_codex":0.009764874,"about_ca_topic_score_gemma":0.013116697,"teacher_disagreement_score":0.0122037,"about_ca_system_score_codex":0.0010168677,"about_ca_system_score_gemma":0.0018029115,"threshold_uncertainty_score":0.06454009},"labels":[],"label_agreement":null},{"id":"W3030521948","doi":"","title":"HardEval: Focusing on Challenging Tokens to Assess Robustness of NER.","year":2020,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Institut de Valorisation des Données","keywords":"Robustness (evolution); Computer science; Ambiguity; Exploit; Artificial intelligence; Machine learning; Natural language processing; Programming language; Computer security","score_opus":0.1133778731560633,"score_gpt":0.2902726815032783,"score_spread":0.176894808347215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3030521948","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2792453,0.006067986,0.63672,0.0017191331,0.0015610502,0.001842288,0.014625073,0.025873931,0.032345187],"genre_scores_gemma":[0.72853535,0.0007565034,0.22131328,0.0008828898,0.00033487205,0.0013600676,0.034336098,0.0045255413,0.007955425],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9827372,0.00825299,0.001463578,0.003379206,0.0036212727,0.0005457695],"domain_scores_gemma":[0.92721164,0.048731387,0.0035028337,0.012043823,0.007209625,0.001300754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016675342,0.0030553814,0.0016997374,0.007067733,0.0020935235,0.0037621418,0.0031862448,0.0031175355,0.0046765367],"category_scores_gemma":[0.07067739,0.0006165184,0.001631938,0.0038141261,0.002384532,0.00847316,0.005765129,0.0036732226,0.003518142],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031421238,0.0014229685,0.082909465,0.005328585,0.0028093397,0.0009735164,0.0024006318,0.23026899,0.039970823,0.016544223,0.07363366,0.54059577],"study_design_scores_gemma":[0.0002552371,0.0020113254,0.050013106,0.00041556213,0.00070121523,0.0020535563,0.0019661156,0.7564154,0.101923876,0.040922254,0.042834315,0.0004879187],"about_ca_topic_score_codex":0.0045604715,"about_ca_topic_score_gemma":0.00788839,"teacher_disagreement_score":0.016675342,"about_ca_system_score_codex":0.0014544519,"about_ca_system_score_gemma":0.0013661738,"threshold_uncertainty_score":0.08818865},"labels":[],"label_agreement":null},{"id":"W3030811146","doi":"","title":"Language Modeling with a General Second-Order RNN.","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Recurrent neural network; Computer science; Multiplicative function; Treebank; Language model; State space; State (computer science); Artificial intelligence; Space (punctuation); Sequence (biology); Artificial neural network; Algorithm; Mathematics; Statistics; Parsing","score_opus":0.030082924266851083,"score_gpt":0.2736927433707653,"score_spread":0.24360981910391424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3030811146","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017295765,0.0008975226,0.9679602,0.0005039775,0.0002601156,0.0002161963,0.0015471446,0.0056049284,0.0057142167],"genre_scores_gemma":[0.5485517,0.00071295194,0.42445138,0.0005517228,0.0002582036,0.0007269519,0.0062896786,0.0012298602,0.017227527],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99756193,0.0013589635,0.00011368353,0.00055492076,0.0002461541,0.00016438843],"domain_scores_gemma":[0.99559116,0.003017338,0.00012361129,0.000411178,0.00069477636,0.00016189019],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037478535,0.0013537433,0.00088107376,0.0011506069,0.00054573163,0.0023175825,0.0019909116,0.0018171361,0.007478758],"category_scores_gemma":[0.012037051,0.00059591996,0.0012204493,0.0010524056,0.0004007799,0.0035728288,0.001491677,0.0028192648,0.005345342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010306039,0.00031895316,0.0025935792,0.00056601997,0.000430009,0.0002837379,0.0004130266,0.36075035,0.013188823,0.021781554,0.017181339,0.58146197],"study_design_scores_gemma":[0.000012908061,0.000039141603,0.00022539415,0.000019988162,0.000030603933,0.000042728243,0.000029395691,0.9884133,0.0023528757,0.0071790935,0.0016424201,0.0000122215615],"about_ca_topic_score_codex":0.013356977,"about_ca_topic_score_gemma":0.017959462,"teacher_disagreement_score":0.013356977,"about_ca_system_score_codex":0.0011226962,"about_ca_system_score_gemma":0.001774111,"threshold_uncertainty_score":0.026558459},"labels":[],"label_agreement":null},{"id":"W3030836112","doi":"","title":"Generalizing Argumentative Link Identification Model by Reducing Dependence on Superficial Cues","year":2020,"lang":"en","type":"article","venue":"The Japanese Society for Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nexen (Canada)","funders":"","keywords":"Argumentative; Identification (biology); Link (geometry); Artificial intelligence; Computer science; Mathematics; Communication; Psychology; Epistemology; Biology; Philosophy; Computer network","score_opus":0.10073973727377551,"score_gpt":0.3239293888584875,"score_spread":0.223189651584712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3030836112","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07523695,0.00027007714,0.89570093,0.0013687203,0.0001342794,0.00014950009,0.00031177874,0.001014276,0.025813568],"genre_scores_gemma":[0.89140105,0.00027359425,0.0991948,0.00031397183,0.00012790543,0.00019525488,0.00056767574,0.00023346918,0.0076922383],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99850225,0.0005764556,0.00008285044,0.00042673605,0.00028344375,0.00012829606],"domain_scores_gemma":[0.99415475,0.0032213759,0.0003688839,0.0010916231,0.0009735221,0.00018981926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025558905,0.0006758631,0.0006603329,0.0021098324,0.00087501074,0.0026056874,0.0017148097,0.0016142152,0.012855267],"category_scores_gemma":[0.016425928,0.0006357609,0.0011614414,0.0012450857,0.0010772198,0.008889662,0.0026504565,0.0025321685,0.0023664627],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063843885,0.00064628286,0.010252284,0.0005252594,0.0003303981,0.000546436,0.0032553812,0.03766971,0.016116967,0.61202496,0.0103748515,0.30761904],"study_design_scores_gemma":[0.00006521523,0.000088115194,0.0033314351,0.00005334054,0.000228418,0.00021692715,0.0005778199,0.55771077,0.0038499746,0.42794845,0.005875865,0.00005370381],"about_ca_topic_score_codex":0.0018181292,"about_ca_topic_score_gemma":0.0016850493,"teacher_disagreement_score":0.012855267,"about_ca_system_score_codex":0.00059945043,"about_ca_system_score_gemma":0.0012222519,"threshold_uncertainty_score":0.04300511},"labels":[],"label_agreement":null},{"id":"W3031043157","doi":"","title":"WEXEA: Wikipedia EXhaustive Entity Annotation","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Annotation; Hyperlink; Information retrieval; Entity linking; Named-entity recognition; Relationship extraction; Natural language processing; Task (project management); Publication; Information extraction; Named entity; Relation (database); Artificial intelligence; World Wide Web; Web page; Knowledge base; Database","score_opus":0.033941972694160195,"score_gpt":0.2878343625034358,"score_spread":0.2538923898092756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3031043157","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12378736,0.006735745,0.13695626,0.0022883858,0.003138864,0.0057238187,0.47861922,0.198221,0.04452935],"genre_scores_gemma":[0.113951795,0.0010466365,0.13273951,0.0007503279,0.00023693552,0.0031164684,0.7258298,0.0084160585,0.013912545],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98602146,0.005578929,0.0016300832,0.0026508265,0.0032578853,0.0008607636],"domain_scores_gemma":[0.96798193,0.014577806,0.00095654605,0.007211316,0.0073012803,0.0019710786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009988223,0.0036354358,0.002159588,0.01348401,0.0038636462,0.0045115845,0.003977229,0.0029544048,0.025079422],"category_scores_gemma":[0.035603423,0.0012709717,0.0019265752,0.0069316532,0.0011989367,0.00945017,0.008674054,0.0028449611,0.017614784],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026461824,0.0018446546,0.010073547,0.005221295,0.0013306037,0.00085056655,0.0013079469,0.0065058754,0.01190307,0.0051574344,0.6825647,0.27059412],"study_design_scores_gemma":[0.002505266,0.0021884916,0.034679085,0.0022721593,0.0022424713,0.0028547859,0.004811017,0.18303807,0.079982124,0.022038521,0.6624891,0.00089895906],"about_ca_topic_score_codex":0.03309637,"about_ca_topic_score_gemma":0.043944776,"teacher_disagreement_score":0.03309637,"about_ca_system_score_codex":0.0014597713,"about_ca_system_score_gemma":0.0052247955,"threshold_uncertainty_score":0.08389902},"labels":[],"label_agreement":null},{"id":"W3032097118","doi":"","title":"Multilingual Corpus Creation for Multilingual Semantic Similarity Task","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Western University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Task (project management); Semantic similarity; Similarity (geometry); Sentence; Focus (optics); Information retrieval","score_opus":0.04671167661381604,"score_gpt":0.32804333242450934,"score_spread":0.2813316558106933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3032097118","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.35983387,0.003532889,0.4620368,0.0016624393,0.0028510084,0.0042803804,0.06052716,0.04745438,0.05782105],"genre_scores_gemma":[0.4802123,0.00089642114,0.35861674,0.0004036339,0.00046344512,0.004078203,0.13491991,0.0045257583,0.015883502],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99586403,0.001499945,0.0005469307,0.0010226321,0.0007505556,0.00031594414],"domain_scores_gemma":[0.9937643,0.0023132055,0.00016183773,0.0010248956,0.0022586256,0.00047712232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034991994,0.0018633584,0.0014283892,0.0063913604,0.003177388,0.0030685237,0.0016775157,0.0015652209,0.021942575],"category_scores_gemma":[0.011471543,0.0006827826,0.0014150613,0.0038132095,0.000722763,0.0056813043,0.0056137377,0.0023508451,0.0104427505],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023537402,0.0016432743,0.009811264,0.003140762,0.000514995,0.0018910714,0.0043509165,0.0079651065,0.078197636,0.02039156,0.183147,0.6865927],"study_design_scores_gemma":[0.0011729263,0.0015919978,0.024218384,0.00068477704,0.0014881822,0.0050814375,0.012245944,0.32118014,0.23382759,0.031117897,0.36673254,0.00065813307],"about_ca_topic_score_codex":0.008061717,"about_ca_topic_score_gemma":0.008850508,"teacher_disagreement_score":0.021942575,"about_ca_system_score_codex":0.0011205545,"about_ca_system_score_gemma":0.0038726877,"threshold_uncertainty_score":0.073405266},"labels":[],"label_agreement":null},{"id":"W3032420555","doi":"","title":"Cooking Up a Neural-based Model for Recipe Classification","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Bank of Canada; Université du Québec à Montréal; Polytechnique Montréal; National Bank of Canada","funders":"","keywords":"Embedding; Computer science; Task (project management); Artificial intelligence; Layer (electronics); Macro; Artificial neural network; Recipe; Natural language processing; Language model; Deep learning; State (computer science); Machine learning; Pattern recognition (psychology); Algorithm; Engineering","score_opus":0.1262190574827426,"score_gpt":0.33255425756928014,"score_spread":0.20633520008653755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3032420555","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12604907,0.0014224768,0.85286915,0.002731309,0.00054282474,0.00029316128,0.0014398124,0.00626941,0.0083828],"genre_scores_gemma":[0.8122604,0.000601353,0.1665064,0.0006877583,0.0002814848,0.00028451695,0.0024697217,0.0003468558,0.016561573],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99935263,0.00020777089,0.00004341989,0.00021234174,0.00010483576,0.000079018646],"domain_scores_gemma":[0.9986155,0.00069225486,0.000041972424,0.00015172083,0.0004315446,0.000066983266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019033289,0.0009264647,0.00083719013,0.0010558597,0.0006376417,0.0017454786,0.0014891529,0.001941561,0.005485358],"category_scores_gemma":[0.004249679,0.00055980054,0.0011008794,0.00088414794,0.000401279,0.0028936656,0.00094042806,0.0028414435,0.0026187405],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092731393,0.0007059289,0.005541811,0.00019267191,0.00044041252,0.00012783178,0.00016874631,0.3341322,0.014720538,0.009172742,0.016085638,0.61778414],"study_design_scores_gemma":[0.000011092352,0.00003962978,0.00034505484,0.000009184121,0.000037286223,0.000016459197,0.000013457865,0.9937894,0.0023219467,0.0028702596,0.00053505506,0.000011150282],"about_ca_topic_score_codex":0.020496408,"about_ca_topic_score_gemma":0.026765937,"teacher_disagreement_score":0.020496408,"about_ca_system_score_codex":0.0014638801,"about_ca_system_score_gemma":0.0011444511,"threshold_uncertainty_score":0.0407542},"labels":[],"label_agreement":null},{"id":"W3033434706","doi":"10.48550/arxiv.2006.03535","title":"CoCon: A Self-Supervised Approach for Controlled Text Generation","year":2020,"lang":"en","type":"preprint","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Natural language processing; Phrase; Artificial intelligence; Transformer; Text generation; Control (management); Language model; Content (measure theory); Mathematics","score_opus":0.03053735471752665,"score_gpt":0.24449856389567617,"score_spread":0.21396120917814954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3033434706","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008819058,0.00011936976,0.9811631,0.000108586086,0.000064956876,0.00012322418,0.00015589494,0.008135115,0.0013107155],"genre_scores_gemma":[0.34628317,0.00015823828,0.6374934,0.000462761,0.00014162666,0.00075823884,0.0015097064,0.0022808695,0.010911983],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99912816,0.00026495283,0.000032424006,0.0003194716,0.00020390959,0.000051080373],"domain_scores_gemma":[0.99774253,0.0013932808,0.0001387837,0.0003275867,0.0002976576,0.00010011871],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012432999,0.000893309,0.00059293484,0.00047452285,0.00035042758,0.0006416203,0.0019542305,0.00093904743,0.0039874366],"category_scores_gemma":[0.004576166,0.00042129247,0.0005328427,0.0003364153,0.0007704565,0.0012527538,0.0013848197,0.0014709864,0.001510586],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006910744,0.000448769,0.0013022678,0.000395335,0.00014164478,0.00033174196,0.00047254076,0.23070548,0.09053544,0.021798639,0.02399323,0.62918377],"study_design_scores_gemma":[0.000023182294,0.000039837505,0.00008011422,0.0000042616466,0.0000056490658,0.000028788112,0.000009649442,0.983624,0.010451347,0.0036382908,0.0020878376,0.000007019777],"about_ca_topic_score_codex":0.0023764963,"about_ca_topic_score_gemma":0.0051847943,"teacher_disagreement_score":0.0039874366,"about_ca_system_score_codex":0.00073160446,"about_ca_system_score_gemma":0.0008215381,"threshold_uncertainty_score":0.013339281},"labels":[],"label_agreement":null},{"id":"W3033611226","doi":"10.18653/v1/2020.acl-demos.27","title":"NLP Scholar: An Interactive Visual Explorer for Natural Language Processing Literature","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Information retrieval; Visualization; Field (mathematics); Citation; Dashboard; Interactive visualization; Information visualization; Data science; Natural language processing; World Wide Web; Artificial intelligence","score_opus":0.03421481912389524,"score_gpt":0.3467810907776469,"score_spread":0.31256627165375167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3033611226","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012825576,0.002839583,0.14454767,0.0026690606,0.0009265199,0.0011666178,0.4532594,0.349693,0.03207255],"genre_scores_gemma":[0.06370561,0.0032166436,0.443464,0.0012716849,0.000663635,0.006962077,0.40301818,0.058118973,0.019579127],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99866414,0.00038563568,0.00014000846,0.00023292021,0.0004683535,0.00010898381],"domain_scores_gemma":[0.99306136,0.004738441,0.0003416734,0.0007210298,0.0005863989,0.0005511677],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.003324407,0.0019143323,0.0011320281,0.012190503,0.0008939472,0.004346732,0.002206135,0.0011058736,0.055359833],"category_scores_gemma":[0.011010896,0.0007110616,0.0015542383,0.008511239,0.0005885886,0.0047667,0.0068131983,0.0022099297,0.016034255],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070293684,0.00019156166,0.0024654765,0.0048935693,0.0003490572,0.0006699461,0.004731501,0.0022741624,0.0078473715,0.02269934,0.8321993,0.12097575],"study_design_scores_gemma":[0.00052274857,0.00008744565,0.0036590646,0.0007142693,0.000076932876,0.00032384903,0.0009440844,0.016174749,0.0035082276,0.036319092,0.9375165,0.00015310614],"about_ca_topic_score_codex":0.0037668897,"about_ca_topic_score_gemma":0.00923625,"teacher_disagreement_score":0.9878095,"about_ca_system_score_codex":0.0009790708,"about_ca_system_score_gemma":0.0020049745,"threshold_uncertainty_score":0.18519711},"labels":[],"label_agreement":null},{"id":"W3033628195","doi":"10.1109/tbdata.2020.2998770","title":"Dynamic Entity-Based Named Entity Recognition Under Unconstrained Tagging Schemes","year":2020,"lang":"en","type":"article","venue":"IEEE Transactions on Big Data","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Named-entity recognition; Natural language processing; Sentence; Artificial intelligence; Word (group theory); Entity linking; Language model; Precision and recall; Named entity; Task (project management); Knowledge base; Linguistics","score_opus":0.15396775093685824,"score_gpt":0.2866244407588694,"score_spread":0.13265668982201115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3033628195","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022297017,0.00044935342,0.9721168,0.00018145116,0.000078425415,0.0000731808,0.00040359012,0.002621623,0.0017785871],"genre_scores_gemma":[0.4020863,0.0010658387,0.57872856,0.00040398445,0.0001171558,0.00022871663,0.006466553,0.00033089795,0.01057207],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99779403,0.0005097963,0.0002137182,0.0010212352,0.0003226384,0.0001385832],"domain_scores_gemma":[0.9965301,0.0010559794,0.00033857967,0.0014086957,0.00058785063,0.00007884087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025215717,0.0011937324,0.0010258075,0.0014394568,0.0007237849,0.001801053,0.0022127735,0.0016776618,0.0014308508],"category_scores_gemma":[0.0056272894,0.0005242612,0.0012396672,0.0019563604,0.0008225462,0.009331742,0.0024102277,0.0018675979,0.0022774132],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047074963,0.00027470972,0.005512995,0.0003539926,0.000259734,0.0010689896,0.000722148,0.30737412,0.04089523,0.04587005,0.010546132,0.58665115],"study_design_scores_gemma":[0.000015038031,0.00007374464,0.001211498,0.000031599066,0.00007658976,0.0003326508,0.00010105183,0.948422,0.020062424,0.02021273,0.009387717,0.000072963936],"about_ca_topic_score_codex":0.004072802,"about_ca_topic_score_gemma":0.0052909893,"teacher_disagreement_score":0.004072802,"about_ca_system_score_codex":0.00069233385,"about_ca_system_score_gemma":0.00092821725,"threshold_uncertainty_score":0.013335526},"labels":[],"label_agreement":null},{"id":"W3033644511","doi":"10.1007/978-3-030-49663-0_10","title":"SHAPed Automated Essay Scoring: Explaining Writing Features’ Contributions to English Writing Organization","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence","score_opus":0.015752658711773327,"score_gpt":0.25879612960165854,"score_spread":0.2430434708898852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3033644511","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6618266,0.0007235493,0.2528069,0.0014855814,0.00028436867,0.00029773416,0.0023387838,0.0033699854,0.07686655],"genre_scores_gemma":[0.9486981,0.00007206349,0.043274816,0.00005031347,0.000041791962,0.00007136645,0.00085948565,0.00031952804,0.006612421],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983203,0.0009023376,0.00009465146,0.00030183853,0.00025965847,0.00012120015],"domain_scores_gemma":[0.9754318,0.01726976,0.001923327,0.0015261861,0.00344215,0.00040679786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003130471,0.0005616144,0.0002647707,0.0015557151,0.0008360883,0.0031793327,0.0009038743,0.00069438457,0.008376143],"category_scores_gemma":[0.03625272,0.00032812377,0.00036807486,0.0015145203,0.00064002624,0.0025017513,0.001341938,0.0010418013,0.0022416934],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009768879,0.00033653763,0.14256155,0.00044851,0.00010266276,0.0008049247,0.013395265,0.012387136,0.017299794,0.04097802,0.02824704,0.74246174],"study_design_scores_gemma":[0.00015383403,0.00044980246,0.2359303,0.0005840131,0.0002211594,0.0013637671,0.009420197,0.52721643,0.03062943,0.13929968,0.054531623,0.00019982108],"about_ca_topic_score_codex":0.0022788176,"about_ca_topic_score_gemma":0.0036616153,"teacher_disagreement_score":0.008376143,"about_ca_system_score_codex":0.0007811539,"about_ca_system_score_gemma":0.0007072766,"threshold_uncertainty_score":0.028021038},"labels":[],"label_agreement":null},{"id":"W3034297959","doi":"","title":"Time-aware Large Kernel Convolutions","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Kernel (algebra); Softmax function; Sequence (biology); Automatic summarization; Convolution (computer science); Algorithm; Computational complexity theory; Artificial intelligence; Theoretical computer science; Deep learning; Mathematics; Artificial neural network; Discrete mathematics","score_opus":0.028168816217734004,"score_gpt":0.23736464375740945,"score_spread":0.20919582753967544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034297959","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027642218,0.00035555058,0.96723425,0.00014633966,0.000054622305,0.000026260359,0.00015035704,0.003279069,0.001111358],"genre_scores_gemma":[0.65852344,0.000430336,0.33130386,0.00020916853,0.000113844275,0.00011029263,0.0009277492,0.0005953312,0.0077860313],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948514,0.000102361555,0.000037960417,0.0001765176,0.0001179521,0.00007999639],"domain_scores_gemma":[0.99901605,0.00043547875,0.0000873019,0.0002398728,0.00015910849,0.000062195686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00080178515,0.00081253075,0.0008704591,0.00061921397,0.0004102167,0.0007759938,0.0012753292,0.00082754315,0.003350844],"category_scores_gemma":[0.003339139,0.00034561378,0.00095836766,0.0009332479,0.0005202663,0.00249359,0.0011462943,0.0011955884,0.0015890418],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005693858,0.0001524464,0.0019285049,0.00019724344,0.00012042827,0.00019540149,0.0002902602,0.29517913,0.04755221,0.03088731,0.007095311,0.6158324],"study_design_scores_gemma":[0.0000052922287,0.000020064665,0.00023100246,0.0000036349863,0.0000116394685,0.00004608199,0.000014656723,0.98316044,0.008061414,0.007307934,0.0011293372,0.000008430127],"about_ca_topic_score_codex":0.0059845266,"about_ca_topic_score_gemma":0.008675354,"teacher_disagreement_score":0.0059845266,"about_ca_system_score_codex":0.0010767515,"about_ca_system_score_gemma":0.0010456698,"threshold_uncertainty_score":0.011899352},"labels":[],"label_agreement":null},{"id":"W3034480528","doi":"10.18653/v1/2020.acl-main.514","title":"A Comprehensive Analysis of Preprocessing for Word Representation Learning in Affective Tasks","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Preprocessor; Computer science; Sentiment analysis; Artificial intelligence; Natural language processing; Task (project management); Word (group theory); Data pre-processing; Representation (politics); Task analysis; Support vector machine; Identification (biology); Machine learning; Linguistics; Engineering","score_opus":0.06387737328250263,"score_gpt":0.33081073336377076,"score_spread":0.2669333600812681,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034480528","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07467667,0.0055605154,0.90568125,0.0013944759,0.0004978292,0.0005135027,0.0017286908,0.0039926944,0.00595443],"genre_scores_gemma":[0.3860445,0.004893622,0.58525115,0.0007761514,0.00038729963,0.0013658232,0.011788104,0.0011104668,0.008382871],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983719,0.00059538876,0.0001523122,0.00035424205,0.00036811663,0.00015798604],"domain_scores_gemma":[0.99148506,0.0053256336,0.00036056398,0.0011338273,0.0015495033,0.00014533647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027511637,0.001605485,0.000745815,0.0014347476,0.0007541634,0.0017754162,0.0010980666,0.00097293966,0.006234732],"category_scores_gemma":[0.017430875,0.00044514582,0.0012824483,0.0017156074,0.00046783144,0.003485817,0.0010855937,0.0028133404,0.004796246],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051391334,0.0004702546,0.003043197,0.00070416,0.00013535508,0.000097333745,0.00031390204,0.013912312,0.051066805,0.006590722,0.010693611,0.9124584],"study_design_scores_gemma":[0.00011723976,0.0016081475,0.02733949,0.0004701659,0.0005276134,0.0005152851,0.0010071374,0.67300624,0.18327829,0.053958565,0.057989832,0.00018197195],"about_ca_topic_score_codex":0.0019253478,"about_ca_topic_score_gemma":0.0039312057,"teacher_disagreement_score":0.006234732,"about_ca_system_score_codex":0.00064533454,"about_ca_system_score_gemma":0.0017022607,"threshold_uncertainty_score":0.020857275},"labels":[],"label_agreement":null},{"id":"W3034545352","doi":"10.24963/ijcai.2020/535","title":"End-to-End Transition-Based Online Dialogue Disentanglement","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Session (web analytics); Coherence (philosophical gambling strategy); Scripting language; Semantics (computer science); Cluster analysis; Construct (python library); Matching (statistics); Artificial intelligence; Information retrieval; Data mining; World Wide Web","score_opus":0.037211585609901954,"score_gpt":0.25358572652846895,"score_spread":0.216374140918567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034545352","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.067629226,0.0014381409,0.9040436,0.00059826096,0.00012805132,0.0005111322,0.0043108817,0.018382529,0.002958091],"genre_scores_gemma":[0.6139973,0.00053069793,0.36167878,0.00031028167,0.00012353055,0.0006780956,0.014423714,0.00068425527,0.007573268],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977998,0.00058421993,0.00015287349,0.0009814182,0.00030453948,0.0001771886],"domain_scores_gemma":[0.9976393,0.0011496918,0.00019981487,0.00042604393,0.00037743367,0.00020775975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019316017,0.0021914493,0.0015522017,0.001494453,0.0007036425,0.0020141092,0.0028040945,0.0015268698,0.0026208502],"category_scores_gemma":[0.006952547,0.0006026266,0.0015567568,0.001213825,0.0006238216,0.0048024612,0.0023502132,0.0033657346,0.0023331596],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002817171,0.0014568748,0.016298598,0.0010930573,0.00038802729,0.0005173179,0.0026718238,0.17814377,0.021673633,0.013327268,0.030062657,0.7315498],"study_design_scores_gemma":[0.000035463443,0.00009290265,0.0013594513,0.000022729655,0.000049497405,0.000100062476,0.00016706633,0.97960496,0.004757854,0.008567752,0.005212356,0.000029956596],"about_ca_topic_score_codex":0.008858539,"about_ca_topic_score_gemma":0.012937456,"teacher_disagreement_score":0.008858539,"about_ca_system_score_codex":0.0009943277,"about_ca_system_score_gemma":0.001466131,"threshold_uncertainty_score":0.017613947},"labels":[],"label_agreement":null},{"id":"W3034552719","doi":"10.18653/v1/2020.acl-main.591","title":"Exploiting Syntactic Structure for Better Language Modeling: A Syntactic Distance Approach","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Canadian Institute for Advanced Research; Université de Montréal","funders":"National Natural Science Foundation of China; Westlake University","keywords":"Treebank; Computer science; Perplexity; Parsing; Natural language processing; Artificial intelligence; Language model; Task (project management); Representation (politics); Parse tree; Syntactic structure; Syntax","score_opus":0.05290269024819936,"score_gpt":0.27581019712734883,"score_spread":0.22290750687914948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034552719","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052618813,0.00044220386,0.9423294,0.0007665807,0.000051788356,0.00004522751,0.00030810453,0.0011531359,0.0022847515],"genre_scores_gemma":[0.683215,0.0005706561,0.31035033,0.00028802178,0.00018383228,0.00013210115,0.0015756673,0.00053556106,0.0031488803],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989538,0.00045603747,0.000053508644,0.0003007103,0.00017182982,0.00006406627],"domain_scores_gemma":[0.99738675,0.0014507807,0.00023067734,0.0005237832,0.0002992158,0.00010873858],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016984062,0.0011812167,0.00088361185,0.0015084171,0.0006581393,0.0015606916,0.0016783224,0.0013358016,0.0022296729],"category_scores_gemma":[0.0049221697,0.0005158578,0.0011601107,0.0016553364,0.00082138315,0.004405997,0.0025963637,0.0026155133,0.0013305176],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003056567,0.0003441041,0.0056908745,0.00031653198,0.00035680062,0.00024351908,0.00064754166,0.33141747,0.031182373,0.083882116,0.0065997643,0.53901327],"study_design_scores_gemma":[0.000011363032,0.00004425799,0.00058188406,0.000011146047,0.000036500398,0.000055639797,0.00003272611,0.9383583,0.0034464675,0.055849113,0.0015503389,0.000022302163],"about_ca_topic_score_codex":0.0014040292,"about_ca_topic_score_gemma":0.00239868,"teacher_disagreement_score":0.0022296729,"about_ca_system_score_codex":0.000879763,"about_ca_system_score_gemma":0.0011066741,"threshold_uncertainty_score":0.008982122},"labels":[],"label_agreement":null},{"id":"W3034760557","doi":"10.18653/v1/2020.acl-main.140","title":"Probing Linguistic Features of Sentence-Level Representations in Relation Extraction","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Wirtschaft und Energie; Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Computer science; Natural language processing; Task (project management); Sentence; Encoder; Artificial intelligence; Feature (linguistics); Relation (database); Focus (optics); ENCODE; Architecture; Contrast (vision); Relationship extraction; Linguistics; Information extraction","score_opus":0.08188147680862856,"score_gpt":0.31426508817764426,"score_spread":0.2323836113690157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034760557","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6182117,0.003899127,0.3575824,0.0011842898,0.0002030075,0.00017034616,0.0037197073,0.009875725,0.0051537384],"genre_scores_gemma":[0.90997106,0.00064189546,0.08094632,0.00017246224,0.00006322919,0.000094661176,0.0058700754,0.00023265502,0.0020076374],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999161,0.00038579074,0.000062023806,0.00022639643,0.00009324116,0.00007149027],"domain_scores_gemma":[0.9963313,0.0025092398,0.00023023816,0.0006105532,0.00025741398,0.0000614486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014859636,0.0009559482,0.00047153473,0.0007820488,0.00029447523,0.0010016012,0.0008494775,0.0008249654,0.0023041384],"category_scores_gemma":[0.008342253,0.00038247017,0.0006694544,0.0009261267,0.00043762405,0.004618178,0.0010890996,0.0017114312,0.0011109217],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008194603,0.00038279564,0.012648356,0.0009658927,0.00020089974,0.00030598167,0.0009731011,0.054508775,0.081187025,0.00699194,0.012948728,0.828067],"study_design_scores_gemma":[0.00007049928,0.0004568047,0.012864416,0.000119553,0.00022930752,0.0004192525,0.00045606334,0.88178855,0.0669183,0.026674883,0.009930783,0.00007154252],"about_ca_topic_score_codex":0.0015117949,"about_ca_topic_score_gemma":0.0032892537,"teacher_disagreement_score":0.0023041384,"about_ca_system_score_codex":0.0005154386,"about_ca_system_score_gemma":0.00051015336,"threshold_uncertainty_score":0.007858574},"labels":[],"label_agreement":null},{"id":"W3034865274","doi":"10.24963/ijcai.2020/555","title":"Teacher-Student Networks with Multiple Decoders for Solving Math Word Problem","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"National Natural Science Foundation of China; National Research Foundation","keywords":"Computer science; Margin (machine learning); Generalization; Word (group theory); Visualization; Decoding methods; Theoretical computer science; Artificial intelligence; Machine learning; Algorithm; Mathematics","score_opus":0.026808822564222148,"score_gpt":0.24564886574334815,"score_spread":0.218840043179126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034865274","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22282895,0.0011970813,0.75672054,0.0021371115,0.00019335034,0.00023290853,0.00065728364,0.0036781454,0.012354607],"genre_scores_gemma":[0.8274302,0.0002952509,0.16043273,0.0004937671,0.00010494426,0.0002857108,0.0015734136,0.00021991775,0.0091641145],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994947,0.00019535104,0.000021648655,0.00017796042,0.00005583498,0.00005446661],"domain_scores_gemma":[0.9987081,0.00082879025,0.0000726457,0.00014080494,0.0001652777,0.000084379455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011269029,0.0011729235,0.0006500234,0.0005413549,0.0005824771,0.0006894473,0.0014610682,0.0016972303,0.004071419],"category_scores_gemma":[0.0051328787,0.00045272108,0.0006985921,0.0004883217,0.0006294642,0.0024988805,0.0015436697,0.002471271,0.001156332],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006114167,0.0003445706,0.004799159,0.00024949855,0.00010710499,0.00026200365,0.0003943048,0.62728846,0.007198429,0.014378278,0.008706624,0.33566013],"study_design_scores_gemma":[0.00002086917,0.000042426305,0.00011684811,0.000006484184,0.000012238299,0.000022737706,0.000030215759,0.990973,0.0016896763,0.0063664746,0.0007143201,0.0000047789945],"about_ca_topic_score_codex":0.0056902813,"about_ca_topic_score_gemma":0.01429978,"teacher_disagreement_score":0.0056902813,"about_ca_system_score_codex":0.0011238015,"about_ca_system_score_gemma":0.0013269958,"threshold_uncertainty_score":0.013620257},"labels":[],"label_agreement":null},{"id":"W3034890132","doi":"","title":"Few-shot Relation Extraction via Bayesian Meta-learning on Task Graphs","year":2020,"lang":"en","type":"article","venue":"International Conference on Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; Université de Montréal","funders":"","keywords":"Relationship extraction; Computer science; Leverage (statistics); Parameterized complexity; Relation (database); Artificial intelligence; Bayesian probability; Graph; Machine learning; Theoretical computer science; Algorithm; Data mining","score_opus":0.08655251143981235,"score_gpt":0.3171795000600261,"score_spread":0.23062698862021375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034890132","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014587517,0.0009291585,0.9793158,0.0003615677,0.000053359447,0.00009749582,0.0005715886,0.002909576,0.0011739415],"genre_scores_gemma":[0.4941717,0.00091387104,0.4899543,0.00083444524,0.00020394701,0.00043464525,0.006250672,0.0007967323,0.0064397943],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99790525,0.0005898841,0.00010305869,0.00089625997,0.0003636877,0.00014190667],"domain_scores_gemma":[0.99690825,0.001813338,0.00024407703,0.0006255217,0.00028135994,0.00012739308],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018758249,0.0022668843,0.002221677,0.002818285,0.001046458,0.0019114601,0.004636619,0.0025253668,0.0026348887],"category_scores_gemma":[0.007531204,0.0014475674,0.0020688667,0.0028010227,0.0011638869,0.0068242396,0.0027087205,0.0030926277,0.0018100861],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004982962,0.00048497302,0.003084904,0.0007033962,0.00037487067,0.00072305166,0.00068322406,0.38399237,0.013225666,0.028994104,0.016034987,0.55120015],"study_design_scores_gemma":[0.000014953505,0.00003763705,0.00033205032,0.000024692376,0.00003641905,0.000092148504,0.000042943117,0.94596547,0.0023835422,0.049286593,0.0017631403,0.000020365704],"about_ca_topic_score_codex":0.0047595664,"about_ca_topic_score_gemma":0.01070057,"teacher_disagreement_score":0.0047595664,"about_ca_system_score_codex":0.0013661771,"about_ca_system_score_gemma":0.0014071743,"threshold_uncertainty_score":0.009920418},"labels":[],"label_agreement":null},{"id":"W3035038672","doi":"10.18653/v1/2020.acl-main.204","title":"DeeBERT: Dynamic Early Exiting for Accelerating BERT Inference","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":301,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Vector Institute","funders":"Vector Institute; Natural Sciences and Engineering Research Council of Canada; Government of Canada; Canadian Institute for Advanced Research","keywords":"Inference; Computer science; Transformer; Language model; Redundancy (engineering); Artificial intelligence; Machine learning; Operating system; Engineering","score_opus":0.07716608598601515,"score_gpt":0.29625498082269247,"score_spread":0.21908889483667732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035038672","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009571702,0.00043159112,0.96442044,0.00031557173,0.00016881834,0.000100387275,0.000520601,0.022730755,0.001740136],"genre_scores_gemma":[0.31486335,0.00052580965,0.6637363,0.0007053777,0.00022505666,0.00031759014,0.004135612,0.0045125773,0.010978328],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990557,0.00021577252,0.00005726882,0.00029504762,0.00024265626,0.00013365652],"domain_scores_gemma":[0.9975017,0.0013702154,0.000118124335,0.0005075997,0.0003532481,0.0001491134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018739759,0.0019993265,0.0012493101,0.0013925756,0.0008672692,0.0018217545,0.0032884956,0.0018383038,0.0107624205],"category_scores_gemma":[0.01125251,0.0012427949,0.0012881721,0.0009307718,0.0008216222,0.0055202027,0.0031426046,0.0046234136,0.0054479362],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011197465,0.00027591953,0.0047401837,0.00040852258,0.00023628687,0.0003643305,0.00051498826,0.19590195,0.020422546,0.035818357,0.04730601,0.6928912],"study_design_scores_gemma":[0.000039954783,0.000035690886,0.00019531783,0.000015523969,0.000021961356,0.000055686174,0.000031989504,0.9761317,0.0059829727,0.0129411165,0.0045294003,0.000018643532],"about_ca_topic_score_codex":0.011949433,"about_ca_topic_score_gemma":0.029682716,"teacher_disagreement_score":0.011949433,"about_ca_system_score_codex":0.0012622621,"about_ca_system_score_gemma":0.0022171976,"threshold_uncertainty_score":0.036003888},"labels":[],"label_agreement":null},{"id":"W3035044482","doi":"10.18653/v1/2020.acl-main.128","title":"Few-shot Slot Tagging with Collapsed Dependency Transfer and Label-enhanced Task-adaptive Projection Network","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":189,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Science North","funders":"National Natural Science Foundation of China","keywords":"Computer science; Artificial intelligence; Shot (pellet); Conditional random field; Dependency (UML); Projection (relational algebra); Similarity (geometry); Task (project management); Semantics (computer science); Transfer of learning; Pattern recognition (psychology); Word (group theory); Representation (politics); Multi-label classification; Natural language processing; Machine learning; Algorithm; Image (mathematics); Mathematics","score_opus":0.03743755894945409,"score_gpt":0.22904664415363782,"score_spread":0.19160908520418374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035044482","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07846897,0.0010164104,0.9112971,0.00050978694,0.00023066314,0.00017673954,0.0005924755,0.004122447,0.0035854413],"genre_scores_gemma":[0.7605639,0.000610311,0.22034945,0.00060943386,0.00021208206,0.00039753478,0.00349478,0.00033200232,0.013430479],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991673,0.00022331576,0.000026270163,0.0003808002,0.00010940047,0.00009295781],"domain_scores_gemma":[0.99864787,0.0006258422,0.000090416084,0.00026909384,0.00024757956,0.000119247175],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011555221,0.00132294,0.0011797459,0.00087454345,0.0008587329,0.00081195217,0.0023841271,0.0016686678,0.0022986566],"category_scores_gemma":[0.0027776489,0.0005568858,0.0009551298,0.0012584667,0.0008046236,0.0043218923,0.0020283503,0.002211639,0.001349005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008064985,0.000843338,0.0070185238,0.00043002528,0.00019256606,0.00072831963,0.0009088663,0.1746259,0.021373117,0.016907072,0.017586945,0.7585789],"study_design_scores_gemma":[0.000013726153,0.000080046884,0.0005452602,0.000015720341,0.000031528405,0.000110693116,0.000058970432,0.9798782,0.0030590745,0.014416686,0.0017673593,0.000022717744],"about_ca_topic_score_codex":0.0049462183,"about_ca_topic_score_gemma":0.009743071,"teacher_disagreement_score":0.0049462183,"about_ca_system_score_codex":0.0008002978,"about_ca_system_score_gemma":0.0013142006,"threshold_uncertainty_score":0.009834886},"labels":[],"label_agreement":null},{"id":"W3035315091","doi":"10.18653/v1/2020.acl-main.355","title":"Two Birds, One Stone: A Simple, Unified Model for Text Generation from Structured and Unstructured Data","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Text generation; Table (database); Unstructured data; Code (set theory); Artificial neural network; Task (project management); Simple (philosophy); Artificial intelligence; Natural language processing; Recurrent neural network; Plain text; Machine learning; Theoretical computer science; Data mining; Programming language; Big data","score_opus":0.12293428738088837,"score_gpt":0.29712727689759044,"score_spread":0.1741929895167021,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035315091","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0072430708,0.00049091835,0.98515356,0.0014508083,0.00018589215,0.00011147281,0.0003714074,0.0019555918,0.0030372161],"genre_scores_gemma":[0.25740808,0.0009109433,0.7176621,0.001399339,0.00036911308,0.0005240904,0.001653821,0.00085208495,0.019220458],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942386,0.00019924878,0.000032036336,0.00020700345,0.00009279007,0.000045022127],"domain_scores_gemma":[0.9989743,0.00047453665,0.000047594232,0.00025412935,0.00015412415,0.000095317504],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018506602,0.00079829775,0.0008611193,0.00095991284,0.00059331866,0.0017693192,0.0020418453,0.0017424081,0.0051182243],"category_scores_gemma":[0.0041917334,0.0006858145,0.0011489716,0.00083006243,0.0012949748,0.0041720686,0.0021963103,0.0026970957,0.002930118],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007767547,0.00026759066,0.0035016055,0.00046555154,0.00029227947,0.0005443103,0.0017025741,0.26879638,0.024389837,0.19658737,0.039049994,0.46362582],"study_design_scores_gemma":[0.000028180299,0.0000650447,0.0002790422,0.000028224998,0.000031255102,0.00009713199,0.000044073437,0.93092394,0.0017549077,0.058445133,0.0082752295,0.000027822096],"about_ca_topic_score_codex":0.006751421,"about_ca_topic_score_gemma":0.010344196,"teacher_disagreement_score":0.006751421,"about_ca_system_score_codex":0.000774417,"about_ca_system_score_gemma":0.0010857783,"threshold_uncertainty_score":0.01712215},"labels":[],"label_agreement":null},{"id":"W3035321115","doi":"10.18653/v1/2020.bionlp-1.19","title":"Extensive Error Analysis and a Learning-Based Evaluation of Medical Entity Recognition Systems to Approximate User Experience","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Classifier (UML); Artificial intelligence; Machine learning; Judgement; Metric (unit); Test set; Medical diagnosis; Natural language processing; Gold standard (test); Information retrieval; Pattern recognition (psychology); Data mining; Mathematics","score_opus":0.08982900735465392,"score_gpt":0.3426717351197669,"score_spread":0.25284272776511296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035321115","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.866016,0.0025833936,0.11896138,0.00035515754,0.0003736332,0.00042846942,0.002866106,0.0047232923,0.003692656],"genre_scores_gemma":[0.9433745,0.00030375458,0.04765461,0.00013748418,0.000060872357,0.00021500554,0.005478891,0.00029911916,0.0024758142],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.984268,0.0066967844,0.002266511,0.0025837102,0.003641378,0.0005436747],"domain_scores_gemma":[0.94769084,0.036525473,0.0026623711,0.0050229146,0.0072321794,0.0008661306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013542241,0.0018205531,0.0010983862,0.0022354852,0.0006373756,0.0014425239,0.0011630689,0.0017128909,0.0016181567],"category_scores_gemma":[0.049630288,0.0002754637,0.0010003714,0.0015409346,0.0008618836,0.0023337766,0.001868255,0.0011888132,0.0009252636],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005951827,0.0019494669,0.0971714,0.0030195576,0.0019633106,0.001261669,0.0031627857,0.20940241,0.04857032,0.00212871,0.01974757,0.60567105],"study_design_scores_gemma":[0.00020605151,0.0056738034,0.10888149,0.0002838125,0.00044166791,0.001911833,0.001257722,0.76267874,0.10518306,0.0024512652,0.010675381,0.00035514543],"about_ca_topic_score_codex":0.0038743452,"about_ca_topic_score_gemma":0.0047339373,"teacher_disagreement_score":0.013542241,"about_ca_system_score_codex":0.001207667,"about_ca_system_score_gemma":0.00048928964,"threshold_uncertainty_score":0.071619034},"labels":[],"label_agreement":null},{"id":"W3035455711","doi":"10.1007/978-3-030-58323-1_16","title":"Evaluating a Multi-sense Definition Generation Model for Multiple Languages","year":2020,"lang":"en","type":"preprint","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Polysemy; Computer science; Contrast (vision); Word (group theory); Context (archaeology); Natural language processing; Artificial intelligence; Language model; Word-sense disambiguation; Linguistics","score_opus":0.21004608936275568,"score_gpt":0.3725867930988983,"score_spread":0.16254070373614263,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035455711","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17128275,0.00055618153,0.813751,0.0012478976,0.00024218652,0.0004809709,0.0011106824,0.008386682,0.0029416534],"genre_scores_gemma":[0.6918979,0.00015914567,0.30105314,0.00033321753,0.00009620911,0.0002904776,0.0030983745,0.00089118036,0.0021802885],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98902404,0.0053618825,0.0007098786,0.0019143439,0.0025097467,0.00048016803],"domain_scores_gemma":[0.9575825,0.032718625,0.0011761318,0.0035137094,0.003948344,0.0010607747],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012289228,0.0014010125,0.0016057251,0.0023081012,0.0009370342,0.0043555507,0.0024739667,0.0025777305,0.004100472],"category_scores_gemma":[0.041957986,0.0007511381,0.0022528972,0.0012816247,0.0009168045,0.005743318,0.0034355235,0.0019657472,0.0011971376],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004480844,0.0013769032,0.019845633,0.001239756,0.0010209905,0.0007841214,0.0017428575,0.4631714,0.022898315,0.044108555,0.013768385,0.4255623],"study_design_scores_gemma":[0.000087795495,0.0001665262,0.00049797277,0.000026629992,0.00010909116,0.000066130866,0.00014166605,0.9810115,0.003387636,0.013640728,0.00084385584,0.00002039105],"about_ca_topic_score_codex":0.005156623,"about_ca_topic_score_gemma":0.0065900064,"teacher_disagreement_score":0.012289228,"about_ca_system_score_codex":0.0020366227,"about_ca_system_score_gemma":0.0024233656,"threshold_uncertainty_score":0.06499243},"labels":[],"label_agreement":null},{"id":"W3035498813","doi":"10.18653/v1/2020.acl-main.362","title":"Graph-to-Tree Learning for Solving Math Word Problems","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":130,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"National Natural Science Foundation of China; National Research Foundation","keywords":"Computer science; Benchmark (surveying); Graph; Encoder; Tree (set theory); Theoretical computer science; Artificial intelligence; Machine learning; Algorithm; Mathematics; Combinatorics","score_opus":0.042135381542589154,"score_gpt":0.24177412935767797,"score_spread":0.19963874781508884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035498813","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04001089,0.0016219259,0.9440437,0.0010841638,0.00013704113,0.00018430878,0.0012806029,0.0049150554,0.006722379],"genre_scores_gemma":[0.36524254,0.001515445,0.6169407,0.0006309066,0.00014846878,0.00047824474,0.006416954,0.0005666742,0.008060124],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99970263,0.00009437916,0.000019238249,0.00009664475,0.00006366396,0.000023456761],"domain_scores_gemma":[0.99943966,0.00038065104,0.000043143267,0.00004850982,0.00006514097,0.000022890566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00047159943,0.0010836294,0.0005487867,0.0008492553,0.00036770452,0.0006318731,0.001081131,0.0010313754,0.004905694],"category_scores_gemma":[0.0032818208,0.00027284207,0.00084076996,0.0015731443,0.00054794573,0.0026340932,0.00072999083,0.0020839146,0.0013864746],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011481338,0.00023660103,0.001110411,0.00069262896,0.00007206386,0.00013587321,0.00028324916,0.32950956,0.005258049,0.0579388,0.029799467,0.57484853],"study_design_scores_gemma":[0.000025026511,0.00003505183,0.00011887672,0.000021613163,0.0000114520735,0.0000161004,0.00003540145,0.9280442,0.0012037635,0.06703949,0.0034418055,0.000007149036],"about_ca_topic_score_codex":0.005504417,"about_ca_topic_score_gemma":0.013148487,"teacher_disagreement_score":0.005504417,"about_ca_system_score_codex":0.00098361,"about_ca_system_score_gemma":0.0010610563,"threshold_uncertainty_score":0.016411185},"labels":[],"label_agreement":null},{"id":"W3036469117","doi":"10.1007/978-3-030-51310-8_21","title":"Towards Explainability in Using Deep Learning for the Detection of Anorexia in Social Media","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Anorexia; Artificial intelligence; Human–computer interaction; Multimedia; Medicine","score_opus":0.04121807602272296,"score_gpt":0.27194564052209813,"score_spread":0.23072756449937518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3036469117","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24234627,0.0035961696,0.7409752,0.0038645307,0.00032904488,0.00015176485,0.0016355081,0.0023873583,0.0047141057],"genre_scores_gemma":[0.88973045,0.001042265,0.10122736,0.00051881355,0.00043041885,0.00013826076,0.0025495992,0.00014247303,0.004220401],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993598,0.0002779965,0.000031165066,0.00016880227,0.00007690189,0.00008520247],"domain_scores_gemma":[0.99534976,0.0037396927,0.00024806437,0.0002818561,0.0002563954,0.00012415413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019851623,0.00081149704,0.0005638177,0.0011051375,0.00031425006,0.0014982497,0.00085304363,0.001473896,0.0024836196],"category_scores_gemma":[0.0077150688,0.00035220265,0.00088635075,0.000645346,0.00050383795,0.0021468522,0.0014507164,0.0020653491,0.00074595545],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011017347,0.0006802894,0.045905333,0.0004644025,0.0005370425,0.00043822516,0.0009915703,0.09426763,0.022906078,0.017768182,0.015501573,0.7994379],"study_design_scores_gemma":[0.000026376463,0.00011768121,0.008607811,0.000058170703,0.000072242045,0.00012418331,0.00011079085,0.9626081,0.0032109981,0.023086116,0.0019576338,0.000019848025],"about_ca_topic_score_codex":0.003020763,"about_ca_topic_score_gemma":0.0030798465,"teacher_disagreement_score":0.003020763,"about_ca_system_score_codex":0.00052398816,"about_ca_system_score_gemma":0.0005146347,"threshold_uncertainty_score":0.010498643},"labels":[],"label_agreement":null},{"id":"W3037078219","doi":"","title":"Reliability of Perplexity to Find Number of Latent Topics.","year":2020,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Perplexity; Computer science; Reliability (semiconductor); Natural language processing; Language model","score_opus":0.0205574233174598,"score_gpt":0.24901130416573194,"score_spread":0.22845388084827215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037078219","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5377262,0.0050453073,0.4234024,0.0010409013,0.000560914,0.0004956674,0.015149287,0.0056711473,0.010908155],"genre_scores_gemma":[0.96883553,0.00038156036,0.022526259,0.000078186036,0.00020504749,0.0002444525,0.0062437747,0.0003983666,0.0010867487],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98885846,0.0060198084,0.00075324316,0.0021660193,0.0016523227,0.00055012753],"domain_scores_gemma":[0.9212006,0.06262065,0.0025192024,0.0075577484,0.004972319,0.0011295651],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014004449,0.0014554994,0.0015453205,0.008782848,0.00068041653,0.0023168968,0.0013690112,0.0017228296,0.0033342312],"category_scores_gemma":[0.08247761,0.0005378546,0.0014820689,0.0042922636,0.0008814339,0.0026287094,0.0018178879,0.002629002,0.0031970674],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037334603,0.00046783156,0.3913461,0.001397917,0.0033400077,0.00040642577,0.0023922168,0.09299412,0.01482892,0.008114137,0.03274977,0.448229],"study_design_scores_gemma":[0.00013670266,0.0005850176,0.15307376,0.00030981062,0.000709149,0.00087570806,0.00066724693,0.79914063,0.01255583,0.025867736,0.005851359,0.00022698211],"about_ca_topic_score_codex":0.004016488,"about_ca_topic_score_gemma":0.0027988222,"teacher_disagreement_score":0.014004449,"about_ca_system_score_codex":0.00095135655,"about_ca_system_score_gemma":0.0008265392,"threshold_uncertainty_score":0.07406354},"labels":[],"label_agreement":null},{"id":"W3037566932","doi":"","title":"Augmenting Named Entity Recognition with Commonsense Knowledge","year":2019,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal; Université du Québec","funders":"","keywords":"Commonsense knowledge; Computer science; Natural language processing; Artificial intelligence; Named-entity recognition; Ambiguity; Knowledge base; Natural language understanding; Natural language; Word embedding; Question answering; Embedding; Task (project management)","score_opus":0.019435820131858655,"score_gpt":0.25008630713282126,"score_spread":0.2306504870009626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037566932","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06969673,0.0011588283,0.90508664,0.00068947766,0.00023229928,0.0001312904,0.0013607218,0.014750937,0.006893104],"genre_scores_gemma":[0.6793693,0.0006678288,0.30787987,0.00033030694,0.00012270981,0.00008289336,0.0063067726,0.0003832106,0.0048571546],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987594,0.0003097872,0.000090046924,0.00048682065,0.0002489571,0.000104988976],"domain_scores_gemma":[0.9962058,0.0016359396,0.00023218161,0.0013554269,0.00048640676,0.00008420384],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021089963,0.0008928347,0.0009279219,0.0030153322,0.00067170564,0.0015091124,0.0016784399,0.0013486659,0.0036639345],"category_scores_gemma":[0.0057560806,0.0004773476,0.0010604181,0.0020458146,0.0008017191,0.010584362,0.0032998163,0.0018702648,0.0028922546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023534073,0.00031393295,0.0045415945,0.0004171685,0.00022369865,0.00064114766,0.0005600667,0.046921395,0.024956662,0.011744256,0.011484149,0.8979607],"study_design_scores_gemma":[0.000019832376,0.0001264048,0.0046038474,0.00010188364,0.00016198984,0.0007415318,0.0003309118,0.8833186,0.039865803,0.045203224,0.025430754,0.00009519204],"about_ca_topic_score_codex":0.0027464477,"about_ca_topic_score_gemma":0.0064456863,"teacher_disagreement_score":0.0036639345,"about_ca_system_score_codex":0.00043626598,"about_ca_system_score_gemma":0.00071599387,"threshold_uncertainty_score":0.01225704},"labels":[],"label_agreement":null},{"id":"W3037585139","doi":"10.18653/v1/2020.repl4nlp-1.10","title":"Exploring the Limits of Simple Learners in Knowledge Distillation for Document Classification with DocBERT","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Distillation; Computer science; Task (project management); Baseline (sea); Artificial intelligence; FLOPS; Machine learning; Natural language processing; Parallel computing; Engineering","score_opus":0.22615006821150216,"score_gpt":0.31474570901882165,"score_spread":0.08859564080731949,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037585139","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4450441,0.008473228,0.4999358,0.0033804125,0.00063400535,0.00022924169,0.0013110901,0.021039857,0.019952245],"genre_scores_gemma":[0.8836865,0.000958595,0.10650417,0.000631494,0.00013726932,0.00012852326,0.0013888734,0.00047938846,0.0060852044],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999329,0.00022258043,0.000035852532,0.00015671323,0.00014076287,0.0001151106],"domain_scores_gemma":[0.99697673,0.0020336963,0.00010070858,0.0005112649,0.00025074795,0.00012691446],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024322027,0.0018008668,0.0009202693,0.00083434104,0.0006292588,0.0017938783,0.0018859628,0.0017037449,0.0033206425],"category_scores_gemma":[0.008321455,0.0005437961,0.00043598056,0.00078646664,0.0009823588,0.0071399724,0.0022336727,0.0035993163,0.0016948382],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009588664,0.0005436139,0.0035682316,0.00051122764,0.0002039787,0.00024410144,0.000364868,0.40153414,0.013922828,0.014484676,0.010438565,0.5532249],"study_design_scores_gemma":[0.000031281194,0.00013179613,0.00025146067,0.000030236595,0.000023817753,0.00004498739,0.000065635235,0.9811778,0.0073523545,0.0090394765,0.001831828,0.000019235153],"about_ca_topic_score_codex":0.007445176,"about_ca_topic_score_gemma":0.013582266,"teacher_disagreement_score":0.007445176,"about_ca_system_score_codex":0.0010147507,"about_ca_system_score_gemma":0.0014887502,"threshold_uncertainty_score":0.014803708},"labels":[],"label_agreement":null},{"id":"W3037664376","doi":"10.18653/v1/2020.acl-srw.42","title":"Compositional Generalization by Factorizing Alignment and Translation","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; Université de Montréal","funders":"Office of Naval Research; University of California, Davis","keywords":"Computer science; Generalization; Natural language processing; Artificial intelligence; Task (project management); Machine translation; Translation (biology); Heuristic; Cognitive architecture; Machine learning; Cognition","score_opus":0.04000444592056415,"score_gpt":0.23538642386950895,"score_spread":0.1953819779489448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037664376","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056793876,0.00036891492,0.93212336,0.00038883643,0.00008469193,0.000078335,0.00029660438,0.005320368,0.0045449594],"genre_scores_gemma":[0.69559443,0.00053684175,0.2895617,0.00040098396,0.00021949186,0.00029402252,0.0025516937,0.0013323302,0.009508518],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99913305,0.00021284304,0.000053547807,0.000403477,0.00011488658,0.00008213024],"domain_scores_gemma":[0.9983094,0.0005537259,0.0001529663,0.0007130421,0.00019990848,0.00007097852],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016332594,0.0017293082,0.0010364298,0.0010437549,0.00060901005,0.0011680195,0.0012941856,0.0010115866,0.0057167765],"category_scores_gemma":[0.005630556,0.0005669473,0.0020860536,0.0011705685,0.0011925088,0.0036859124,0.0027901132,0.002171341,0.0033315876],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055674615,0.00032333413,0.0039140657,0.00033677655,0.00027593083,0.00036436837,0.0009408281,0.2897642,0.048860982,0.07744652,0.012813475,0.56440276],"study_design_scores_gemma":[0.000039520244,0.00009012098,0.00046634118,0.00002314359,0.000071814946,0.000107018044,0.00006981924,0.89285713,0.009065266,0.09165054,0.0055306153,0.000028710023],"about_ca_topic_score_codex":0.0027504198,"about_ca_topic_score_gemma":0.0064547625,"teacher_disagreement_score":0.0057167765,"about_ca_system_score_codex":0.0007008123,"about_ca_system_score_gemma":0.0012351084,"threshold_uncertainty_score":0.019124508},"labels":[],"label_agreement":null},{"id":"W3037980084","doi":"10.1109/cist49399.2021.9357170","title":"Leveraging Subword Embeddings for Multinational Address Parsing","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Parsing; Computer science; Python (programming language); Artificial intelligence; Disk formatting; Natural language processing; Artificial neural network; Machine learning; Programming language","score_opus":0.09820144419307504,"score_gpt":0.3210443635176114,"score_spread":0.22284291932453637,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037980084","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17132166,0.0016746903,0.68156046,0.0011541832,0.0012155834,0.00033677687,0.009345814,0.116245694,0.017145082],"genre_scores_gemma":[0.49097395,0.00066234014,0.44277662,0.0007349933,0.00021621895,0.0003160287,0.043018546,0.0037831345,0.017518077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998637,0.00034238803,0.00012602612,0.00054042507,0.00020375763,0.00015032632],"domain_scores_gemma":[0.9980253,0.0006226709,0.00014426654,0.0006801304,0.0004575543,0.000070029375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010441083,0.0026884687,0.0009793582,0.002186654,0.00065343495,0.0020587211,0.0018523583,0.0020359491,0.0075233276],"category_scores_gemma":[0.005327415,0.0007173602,0.001497754,0.0021779248,0.000511897,0.0074782344,0.0026159773,0.0027206698,0.011157515],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034299833,0.0004926946,0.006597823,0.00036236722,0.00024069438,0.00065832224,0.0005122029,0.07101977,0.015786998,0.008411117,0.050825715,0.8447494],"study_design_scores_gemma":[0.00003158687,0.00011360451,0.001397422,0.000062365405,0.000091317306,0.00026175388,0.0003152799,0.9257835,0.02564078,0.021264702,0.02497255,0.000064942345],"about_ca_topic_score_codex":0.005362459,"about_ca_topic_score_gemma":0.009224807,"teacher_disagreement_score":0.0075233276,"about_ca_system_score_codex":0.00080102345,"about_ca_system_score_gemma":0.0011285702,"threshold_uncertainty_score":0.025168061},"labels":[],"label_agreement":null},{"id":"W3038022971","doi":"10.65109/hegb2577","title":"Explainable and Contextual Preferences based Decision Making with Assumption-based Argumentation for Diagnostics and Prognostics of Alzheimer's Disease","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Dialogical self; Argumentation theory; Prognostics; Context (archaeology); Computer science; Construct (python library); Artificial intelligence; Defeasible estate; Machine learning; Preference; Management science; Psychology; Data mining; Epistemology; Social psychology; Mathematics; Engineering","score_opus":0.04925085198054838,"score_gpt":0.28163131929949914,"score_spread":0.23238046731895076,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3038022971","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07069211,0.0006758272,0.91944206,0.003293284,0.0000735064,0.00018409498,0.00032866525,0.0003719736,0.004938353],"genre_scores_gemma":[0.6441482,0.0002603425,0.35366544,0.00022177174,0.00006990094,0.00022361956,0.00040805785,0.0000377635,0.0009650359],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9930698,0.004967698,0.00044952182,0.0006204815,0.0006276668,0.0002648443],"domain_scores_gemma":[0.9720079,0.023910344,0.0014676857,0.00092477497,0.0010995404,0.0005897875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007850145,0.000983723,0.00073596445,0.0016792286,0.0011155604,0.0030550112,0.0020443245,0.0022206507,0.003607202],"category_scores_gemma":[0.029510489,0.00053999975,0.0020093818,0.0011041278,0.0014070737,0.0044643516,0.0027509846,0.0029916214,0.00036482944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069796055,0.00054956495,0.009334576,0.00082347065,0.0004620485,0.00091733027,0.004651365,0.5286439,0.0037491862,0.2930855,0.0038781154,0.15320702],"study_design_scores_gemma":[0.00006363313,0.000060530565,0.0009788461,0.00007865005,0.000082337734,0.00011116904,0.0002829894,0.8342894,0.00082406594,0.1602745,0.0029096873,0.000044273325],"about_ca_topic_score_codex":0.0037193755,"about_ca_topic_score_gemma":0.0041188467,"teacher_disagreement_score":0.007850145,"about_ca_system_score_codex":0.0017332713,"about_ca_system_score_gemma":0.0021308125,"threshold_uncertainty_score":0.041516006},"labels":[],"label_agreement":null},{"id":"W3040446934","doi":"10.1177/0265532220937830","title":"More efficient processes for creating automated essay scoring frameworks: A demonstration of two algorithms","year":2020,"lang":"en","type":"article","venue":"Language Testing","topic":"Topic Modeling","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Rubric; Artificial intelligence; Machine learning; Computer science; Support vector machine; Convolutional neural network; Feature engineering; Deep learning; Artificial neural network; Natural language processing; Meaning (existential); Feature (linguistics); Strengths and weaknesses; F1 score; Algorithm; Mathematics; Mathematics education; Linguistics; Psychology","score_opus":0.0348837319221498,"score_gpt":0.30971148316334696,"score_spread":0.2748277512411972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3040446934","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017618459,0.00014805316,0.9709029,0.00033947593,0.00013898258,0.00035376666,0.00018299396,0.0073189735,0.0029963304],"genre_scores_gemma":[0.1304706,0.000119029675,0.864866,0.0000875427,0.00008666054,0.00035599706,0.00045089802,0.00045328803,0.0031100363],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99120927,0.0032655546,0.00083529373,0.0014749842,0.0028805584,0.00033440947],"domain_scores_gemma":[0.9843061,0.005490905,0.0009894988,0.003672051,0.0048962263,0.0006451946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008945438,0.0013842066,0.0010488599,0.0028678663,0.0006793856,0.0032235489,0.0022656021,0.0014739535,0.005215978],"category_scores_gemma":[0.03084167,0.0006046362,0.0009189309,0.0013767404,0.0008801701,0.0040600374,0.0039327038,0.002276634,0.0033694266],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036456683,0.00046583585,0.0055950866,0.00019982447,0.00009232007,0.00017311797,0.00080955884,0.025253091,0.018724134,0.027876733,0.00933248,0.91111314],"study_design_scores_gemma":[0.00009948329,0.00022080555,0.0040405095,0.00006183771,0.000032698623,0.0003489883,0.00030757955,0.9240695,0.031339303,0.02094692,0.018422658,0.00010967359],"about_ca_topic_score_codex":0.0040064002,"about_ca_topic_score_gemma":0.003722943,"teacher_disagreement_score":0.008945438,"about_ca_system_score_codex":0.0011260927,"about_ca_system_score_gemma":0.001913853,"threshold_uncertainty_score":0.047308564},"labels":[],"label_agreement":null},{"id":"W3042667808","doi":"10.18653/v1/2020.emnlp-demos.7","title":"AdapterHub: A Framework for Adapting Transformers","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research","funders":"Bundesministerium für Bildung und Forschung; Deutsche Forschungsgemeinschaft; Samsung; DeepMind; Samsung Advanced Institute of Technology","keywords":"Computer science; Adapter (computing); Bottleneck; Upload; Scalability; Scripting language; Transformer; Artificial intelligence; Distributed computing; Software engineering; World Wide Web; Embedded system; Programming language; Operating system","score_opus":0.0962412776325239,"score_gpt":0.3059694317501878,"score_spread":0.20972815411766393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3042667808","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002070819,0.00045841347,0.7911213,0.00035715936,0.00024529058,0.00034845702,0.0025027788,0.19659309,0.006302663],"genre_scores_gemma":[0.07467129,0.0018141153,0.8151979,0.0012095681,0.0002884409,0.0014579182,0.017245764,0.07352491,0.01459013],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99550784,0.0013488004,0.00057199947,0.0009280678,0.0012240032,0.00041933908],"domain_scores_gemma":[0.9887715,0.0044000833,0.00029845838,0.005252697,0.00088281406,0.00039439587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007512625,0.0023384942,0.0016376863,0.0029206644,0.0013942915,0.0046247393,0.0061717606,0.0023871684,0.032383814],"category_scores_gemma":[0.023264252,0.003050901,0.003655462,0.002813965,0.0019058434,0.016132766,0.0110674985,0.004446851,0.015622683],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016343766,0.00041756008,0.0038301083,0.0017806647,0.00042845643,0.00083963975,0.0012490377,0.015329501,0.013154677,0.22164641,0.26336154,0.47632802],"study_design_scores_gemma":[0.0004600918,0.00013787612,0.0006221,0.00044421462,0.0002558038,0.0007094814,0.00033996033,0.13029456,0.025625935,0.31185117,0.52904636,0.00021252548],"about_ca_topic_score_codex":0.004758017,"about_ca_topic_score_gemma":0.0066902298,"teacher_disagreement_score":0.032383814,"about_ca_system_score_codex":0.0013583674,"about_ca_system_score_gemma":0.0025100245,"threshold_uncertainty_score":0.10833472},"labels":[],"label_agreement":null},{"id":"W3043078559","doi":"10.3390/computers9030057","title":"ERF: An Empirical Recommender Framework for Ascertaining Appropriate Learning Materials from Stack Overflow Discussions","year":2020,"lang":"en","type":"article","venue":"Computers","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Computer science; Python (programming language); Recommender system; Documentation; World Wide Web; Empirical research; Matching (statistics); Information retrieval; Artificial intelligence; Machine learning; Programming language","score_opus":0.0953893547179278,"score_gpt":0.3337515061552247,"score_spread":0.2383621514372969,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3043078559","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0984103,0.0013215906,0.88592005,0.0008580756,0.00007689852,0.0011129684,0.005436827,0.0033317928,0.0035315729],"genre_scores_gemma":[0.49872386,0.0005492767,0.48459315,0.00022254429,0.00016400845,0.0011236969,0.010379391,0.00018034977,0.0040636733],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99342936,0.00368429,0.0005409968,0.0012360957,0.00089762674,0.00021166493],"domain_scores_gemma":[0.9643583,0.027196052,0.0018536948,0.002329993,0.0036700114,0.00059193064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014429464,0.0014751668,0.0015230378,0.00883244,0.0011960873,0.0022439577,0.0030205431,0.0024995224,0.00534017],"category_scores_gemma":[0.044318985,0.0007519976,0.0012315366,0.004455893,0.00071736873,0.003773708,0.0016046071,0.0018768832,0.0019423473],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001627039,0.0020584958,0.18399607,0.0019134107,0.0010286195,0.00052900845,0.0045533,0.073605195,0.009241971,0.03156964,0.028040467,0.6618367],"study_design_scores_gemma":[0.00013206057,0.00043939156,0.028205639,0.00016558942,0.00018327871,0.00023864828,0.0008068179,0.94003254,0.0028617985,0.01585792,0.010934224,0.00014197043],"about_ca_topic_score_codex":0.019102113,"about_ca_topic_score_gemma":0.03197663,"teacher_disagreement_score":0.019102113,"about_ca_system_score_codex":0.0014337349,"about_ca_system_score_gemma":0.0018328924,"threshold_uncertainty_score":0.07631123},"labels":[],"label_agreement":null},{"id":"W3043222912","doi":"10.3389/frai.2020.00042","title":"Using Topic Modeling Methods for Short-Text Data: A Comparative Analysis","year":2020,"lang":"en","type":"article","venue":"Frontiers in Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":368,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Latent Dirichlet allocation; Topic model; Computer science; Latent semantic analysis; Popularity; Information retrieval; Social media; Matrix decomposition; Context (archaeology); Artificial intelligence; Precision and recall; Natural language processing; Data science; Machine learning; Data mining; World Wide Web","score_opus":0.49520728568712485,"score_gpt":0.4730338297708425,"score_spread":0.022173455916282336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3043222912","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14324172,0.14043039,0.6987153,0.0029353995,0.0011111735,0.000753464,0.004320539,0.0020749944,0.006416977],"genre_scores_gemma":[0.5255121,0.056068674,0.397253,0.00040593185,0.0021109849,0.0010265723,0.015015284,0.0006730448,0.0019344162],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9852224,0.008075548,0.0012041078,0.0019458283,0.003314315,0.00023776722],"domain_scores_gemma":[0.88105035,0.10529968,0.0029554563,0.0035128153,0.0065631256,0.00061869225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021623412,0.0016286179,0.0016706086,0.013627423,0.0008241962,0.003625943,0.0014206757,0.0018171681,0.0012885008],"category_scores_gemma":[0.057171937,0.0004620613,0.0025181489,0.012025652,0.00072484824,0.0076519367,0.0011708066,0.0015688221,0.0010111661],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010888894,0.00040430707,0.048071142,0.005398154,0.0024979194,0.00031302765,0.0025532665,0.035734333,0.003815238,0.011696643,0.012178404,0.8762488],"study_design_scores_gemma":[0.00018904722,0.0010841553,0.09188557,0.0018838771,0.0019724627,0.0017145242,0.0043859114,0.7781143,0.008227656,0.03391255,0.07618051,0.00044944102],"about_ca_topic_score_codex":0.0027954346,"about_ca_topic_score_gemma":0.0025546702,"teacher_disagreement_score":0.021623412,"about_ca_system_score_codex":0.0010839744,"about_ca_system_score_gemma":0.0011227535,"threshold_uncertainty_score":0.114356935},"labels":[],"label_agreement":null},{"id":"W3043396802","doi":"10.1037/xlm0000941","title":"Comparing recollection and nonrecollection memory states for recall of general knowledge: A nontrivial pursuit.","year":2020,"lang":"en","type":"article","venue":"Journal of Experimental Psychology Learning Memory and Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Recall; Cognitive psychology; Psychology","score_opus":0.058088003070488795,"score_gpt":0.3346082704187315,"score_spread":0.2765202673482427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3043396802","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9654598,0.0049801725,0.01216494,0.0008790992,0.00008974767,0.00020462378,0.0013579833,0.00013309297,0.014730441],"genre_scores_gemma":[0.9946516,0.0006753884,0.003143708,0.00010384515,0.000036198628,0.00008444911,0.0005518118,0.00002073169,0.00073227775],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9986475,0.00059692294,0.00021043382,0.00023137989,0.0002537641,0.00005994064],"domain_scores_gemma":[0.9562216,0.031140309,0.0044826106,0.005375197,0.0020200212,0.0007602588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005941413,0.00029136616,0.00039037195,0.0016458248,0.00031240607,0.0022928088,0.0009781445,0.00071949116,0.0063016317],"category_scores_gemma":[0.06376813,0.00016331747,0.0003597529,0.002390283,0.0007932418,0.0053794603,0.0012886929,0.0007619565,0.000941085],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001624398,0.0005033783,0.570954,0.001138786,0.00041635864,0.00020234237,0.009328299,0.00048576322,0.007236919,0.0061630807,0.002604353,0.39934233],"study_design_scores_gemma":[0.00003741526,0.0009783335,0.9651894,0.00036285352,0.00017517901,0.0007986814,0.005803827,0.0030668676,0.005360383,0.014693446,0.0034600622,0.000073504576],"about_ca_topic_score_codex":0.00128177,"about_ca_topic_score_gemma":0.0021292395,"teacher_disagreement_score":0.0063016317,"about_ca_system_score_codex":0.00039939847,"about_ca_system_score_gemma":0.0004292267,"threshold_uncertainty_score":0.031421542},"labels":[],"label_agreement":null},{"id":"W3043507478","doi":"10.4018/978-1-7998-3476-2.ch003","title":"Automated Essay Scoring Using Deep Learning Algorithms","year":2020,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Medical Council of Canada; University of Alberta","funders":"","keywords":"Interpretability; Artificial intelligence; Computer science; Field (mathematics); Machine learning; Deep learning; Data science; Mathematics","score_opus":0.03362719573425966,"score_gpt":0.2634630590236776,"score_spread":0.22983586328941796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3043507478","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022685744,0.0014426879,0.9624101,0.0006017558,0.00019045052,0.00010143256,0.00047637091,0.004121247,0.00797022],"genre_scores_gemma":[0.45829156,0.0021154464,0.5003014,0.00029274457,0.00036808816,0.00032059324,0.003195879,0.00046446334,0.034649674],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986156,0.0005414881,0.00010901542,0.00024410505,0.0003961727,0.0000934915],"domain_scores_gemma":[0.99622947,0.001800608,0.00029215368,0.00039104326,0.0011829771,0.00010374697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021232702,0.0010046173,0.00082191627,0.0017507201,0.00034535263,0.0022584575,0.0014022092,0.0008371031,0.0048510414],"category_scores_gemma":[0.008133397,0.00035155672,0.00056018593,0.0016110842,0.000330968,0.0018344406,0.0015832357,0.0017752436,0.0039638984],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006703224,0.00009153758,0.001708631,0.00014110915,0.000047130878,0.000045143403,0.000121276615,0.09506281,0.0035756782,0.009340408,0.011891917,0.87790734],"study_design_scores_gemma":[0.0000068440118,0.000025218867,0.0007252057,0.00003544079,0.000008839039,0.00002615147,0.00003034431,0.976838,0.0025573815,0.015184764,0.0045493813,0.0000125684155],"about_ca_topic_score_codex":0.002809501,"about_ca_topic_score_gemma":0.004531068,"teacher_disagreement_score":0.0048510414,"about_ca_system_score_codex":0.0010020923,"about_ca_system_score_gemma":0.0008359226,"threshold_uncertainty_score":0.016228378},"labels":[],"label_agreement":null},{"id":"W3044693911","doi":"10.1017/s1351324920000509","title":"Comparison of rule-based and neural network models for negation detection in radiology reports","year":2020,"lang":"en","type":"article","venue":"Natural Language Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Population and Public Health","funders":"Engineering and Physical Sciences Research Council; Medical Research Council","keywords":"Computer science; Artificial intelligence; Natural language processing; Negation; Artificial neural network; Sentence; Syntax; Pipeline (software); Machine learning; Rule-based system; Python (programming language); Information retrieval; Programming language","score_opus":0.014563883407551955,"score_gpt":0.24777112941400728,"score_spread":0.23320724600645532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3044693911","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5326376,0.0058492064,0.43569937,0.0036084526,0.00057265867,0.0005153365,0.0034192207,0.0077259648,0.009972169],"genre_scores_gemma":[0.910329,0.0006360064,0.08349852,0.0005368945,0.000119071025,0.00017253683,0.0024547828,0.00012280505,0.002130312],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989027,0.00036643056,0.00011679511,0.00030582887,0.000228522,0.00007973779],"domain_scores_gemma":[0.9859061,0.011605394,0.00058524875,0.00031284848,0.0014164426,0.00017382925],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004678042,0.0012416691,0.0010504443,0.0016804346,0.000432223,0.0019191223,0.0023856722,0.0015261213,0.0027010208],"category_scores_gemma":[0.0154606495,0.0005410477,0.0011469333,0.0008393092,0.0004920559,0.0022790572,0.0007918771,0.0017399978,0.0009059091],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009978489,0.00044433514,0.009266322,0.00026931326,0.00031936204,0.00023350949,0.00011754007,0.84079385,0.0014755734,0.0020242713,0.003461715,0.14059636],"study_design_scores_gemma":[0.000009984504,0.00002563378,0.0002708659,0.0000123686095,0.000016785634,0.000011251044,0.000005782727,0.998372,0.00028805103,0.0008799936,0.000102265294,0.0000050490685],"about_ca_topic_score_codex":0.018603964,"about_ca_topic_score_gemma":0.013474505,"teacher_disagreement_score":0.018603964,"about_ca_system_score_codex":0.0018728069,"about_ca_system_score_gemma":0.0012761076,"threshold_uncertainty_score":0.0369913},"labels":[],"label_agreement":null},{"id":"W3046434790","doi":"10.1007/978-3-030-41251-7_5","title":"Inferring Systemic Nets with Applications to Islamist Forums","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in social networks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Islam; Framing (construction); Computer science; Representation (politics); Headline; Systemic functional linguistics; Mindset; Perspective (graphical); Abstraction; Artificial intelligence; Epistemology; Political science; Linguistics; History","score_opus":0.017921511045729605,"score_gpt":0.2421219794700011,"score_spread":0.2242004684242715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3046434790","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20710085,0.0011156352,0.77281934,0.0008547678,0.00013191132,0.00017827135,0.0057337056,0.0034974697,0.008568017],"genre_scores_gemma":[0.69526964,0.00049346115,0.28884277,0.00006490001,0.00020533154,0.00013885915,0.009011862,0.00024518959,0.0057280967],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99932265,0.00030310408,0.00003957781,0.00017606576,0.00010090234,0.00005776441],"domain_scores_gemma":[0.99480313,0.004249946,0.00021071792,0.00027700697,0.0002995911,0.00015963691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001607759,0.00063477637,0.00046905025,0.0030279385,0.0009904858,0.0016663483,0.00066610216,0.0009806145,0.0032672056],"category_scores_gemma":[0.00687546,0.0004571505,0.00080687634,0.0027804978,0.00041179862,0.0026684767,0.0013771305,0.0009716956,0.0012330227],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075816986,0.0005554734,0.09053268,0.000534055,0.0003424144,0.0008961915,0.0020834594,0.17403783,0.0067964788,0.07968899,0.027247932,0.61652625],"study_design_scores_gemma":[0.000018904604,0.000026986092,0.0055044857,0.000032372627,0.000047835063,0.000100498226,0.00032693867,0.93531746,0.0012487861,0.051638465,0.0057230797,0.0000141268365],"about_ca_topic_score_codex":0.0060254834,"about_ca_topic_score_gemma":0.018844433,"teacher_disagreement_score":0.0060254834,"about_ca_system_score_codex":0.0006981089,"about_ca_system_score_gemma":0.000586663,"threshold_uncertainty_score":0.011980832},"labels":[],"label_agreement":null},{"id":"W3046929991","doi":"10.2196/19848","title":"Chinese Clinical Named Entity Recognition in Electronic Medical Records: Development of a Lattice Long Short-Term Memory Model With Contextualized Character Representations","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Natural language processing; Conditional random field; Artificial intelligence; Ambiguity; Deep learning; Interpretability; Chinese characters; Character (mathematics); Information retrieval","score_opus":0.057220088336350776,"score_gpt":0.36315011546831555,"score_spread":0.30593002713196477,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3046929991","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09577909,0.0027402642,0.88686913,0.0018620659,0.0003417647,0.00025960247,0.002840614,0.005601142,0.003706331],"genre_scores_gemma":[0.6851465,0.0021613198,0.29809457,0.0010070904,0.00024481877,0.00045779234,0.0066797263,0.0001560745,0.0060521215],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99942243,0.00011144216,0.00006598575,0.00023364145,0.000102486505,0.00006400329],"domain_scores_gemma":[0.99897695,0.00040944997,0.000144376,0.00013533205,0.00027616214,0.00005780994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010706323,0.0007788066,0.000770574,0.0011133399,0.00041133273,0.000924867,0.0019467804,0.0009961695,0.0017799024],"category_scores_gemma":[0.0033382047,0.00032904505,0.0009959514,0.001642657,0.00042817352,0.0027168465,0.00095837895,0.0014913597,0.0010364471],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037026507,0.000293208,0.00876626,0.0002946783,0.00022111362,0.0005307792,0.00036315768,0.27741593,0.007071679,0.0121756345,0.013998513,0.67849874],"study_design_scores_gemma":[0.0000126873465,0.00004556809,0.0006008372,0.00001936075,0.000037856433,0.00007135758,0.00003429565,0.9913523,0.001921384,0.004025403,0.0018613166,0.000017722426],"about_ca_topic_score_codex":0.017367557,"about_ca_topic_score_gemma":0.018938294,"teacher_disagreement_score":0.017367557,"about_ca_system_score_codex":0.0010496436,"about_ca_system_score_gemma":0.0022410385,"threshold_uncertainty_score":0.034532905},"labels":[],"label_agreement":null},{"id":"W3047016605","doi":"10.21437/interspeech.2020-2557","title":"To BERT or not to BERT: Comparing Speech and Language-Based Approaches for Alzheimer’s Disease Detection","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Artificial intelligence; Transformer; Encoder; Language model; Feature (linguistics); Task (project management); Natural language processing; Speech recognition; Multi-task learning; Machine learning; Linguistics","score_opus":0.17572904925378088,"score_gpt":0.31966023758658146,"score_spread":0.14393118833280058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3047016605","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8785842,0.016662942,0.08098623,0.0029828558,0.001297229,0.00034304493,0.0049205716,0.007867825,0.0063550808],"genre_scores_gemma":[0.9422538,0.00154248,0.03823011,0.00056615396,0.00049790886,0.000115023235,0.0130658,0.0003808449,0.0033478811],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99665976,0.0016597722,0.00020993786,0.0008736253,0.0003635852,0.00023335453],"domain_scores_gemma":[0.99139035,0.0060964795,0.0003534275,0.00092709385,0.0008329927,0.00039964952],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077910456,0.0023972017,0.0012687535,0.0026781631,0.0006211142,0.0017519313,0.0014925998,0.0024850967,0.0010509408],"category_scores_gemma":[0.012964106,0.0003380885,0.0013164268,0.0011032547,0.00061786186,0.0030667651,0.0017132562,0.0023835567,0.0017216941],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0074084746,0.0028170287,0.049298376,0.0011519467,0.0018284906,0.00030081594,0.0008000885,0.114867784,0.01326303,0.0021313252,0.031628735,0.77450395],"study_design_scores_gemma":[0.00034169824,0.0019716253,0.025133794,0.000129818,0.00048688933,0.0004009674,0.00077522127,0.945489,0.0108862985,0.008031035,0.0061938968,0.00015968284],"about_ca_topic_score_codex":0.010803406,"about_ca_topic_score_gemma":0.011680893,"teacher_disagreement_score":0.010803406,"about_ca_system_score_codex":0.00094292767,"about_ca_system_score_gemma":0.0011433527,"threshold_uncertainty_score":0.04120344},"labels":[],"label_agreement":null},{"id":"W3061846580","doi":"10.48550/arxiv.2008.07680","title":"An Annotated Corpus of Webtables for Information Extraction Tasks","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Information retrieval; Annotation; Relationship extraction; Paragraph; Information extraction; Natural language processing; Task (project management); Benchmark (surveying); Table (database); Sentence; Artificial intelligence; Context (archaeology); Data mining; World Wide Web","score_opus":0.09069754288968365,"score_gpt":0.2105369483483559,"score_spread":0.11983940545867225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3061846580","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052933775,0.004267576,0.03895419,0.0008855636,0.0006357269,0.00069951103,0.8662207,0.011068838,0.024334121],"genre_scores_gemma":[0.032059416,0.0007985174,0.038054246,0.00017565205,0.00009139688,0.000721174,0.9236726,0.00065192423,0.0037751158],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99766254,0.00049947336,0.00031213855,0.0007125176,0.00068517635,0.0001281762],"domain_scores_gemma":[0.99245656,0.0030690741,0.0005148007,0.0015900813,0.0019970743,0.00037250083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012754775,0.0010427133,0.0007142268,0.007687509,0.0015470616,0.0016567538,0.0016233874,0.0014789334,0.011852674],"category_scores_gemma":[0.010122683,0.00048272638,0.00081562396,0.011358201,0.0006403762,0.0026404762,0.0018208451,0.0011957763,0.010934661],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069712405,0.00050483947,0.01084625,0.0072028213,0.00021464824,0.0022370967,0.0022157356,0.004885963,0.020671563,0.011000571,0.7351073,0.20441604],"study_design_scores_gemma":[0.00010723797,0.000115542585,0.030760003,0.0007124179,0.000119085074,0.0015505493,0.0013381667,0.0090850685,0.016463302,0.006578849,0.93304956,0.00012010162],"about_ca_topic_score_codex":0.011455171,"about_ca_topic_score_gemma":0.024705378,"teacher_disagreement_score":0.011852674,"about_ca_system_score_codex":0.0011025192,"about_ca_system_score_gemma":0.0025252132,"threshold_uncertainty_score":0.039651155},"labels":[],"label_agreement":null},{"id":"W3071625105","doi":"10.18653/v1/2021.eacl-main.228","title":"Do Syntax Trees Help Pre-trained Transformers Extract Information?","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Canadian Institute for Advanced Research; Microsoft Research","keywords":"Computer science; Transformer; ENCODE; Syntax; Abstract syntax tree; Artificial intelligence; Abstract syntax; Natural language processing; Dependency (UML); Dependency graph; Machine learning; Graph; Theoretical computer science; Engineering","score_opus":0.019409812430269693,"score_gpt":0.2625592732915782,"score_spread":0.2431494608613085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3071625105","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16153581,0.0062151905,0.7878903,0.0045434004,0.00062569004,0.0002670418,0.005786999,0.020645712,0.012489873],"genre_scores_gemma":[0.733867,0.003231163,0.23505358,0.0011933331,0.00022175467,0.00017282498,0.017896531,0.0017359266,0.0066279494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99856347,0.0006177736,0.0000857726,0.00045157943,0.00014049711,0.00014088445],"domain_scores_gemma":[0.9907864,0.006812802,0.00030792353,0.0012624931,0.00060860714,0.00022171879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038253432,0.002551239,0.0009844655,0.0017982688,0.00053932203,0.0025110801,0.0019234252,0.0016659058,0.005593807],"category_scores_gemma":[0.018716535,0.0009251142,0.0015450036,0.0016211716,0.0006992535,0.013140964,0.0016795269,0.004065131,0.0064481236],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007131371,0.00041616696,0.017238075,0.0008837597,0.000495029,0.00037031394,0.00069633673,0.072863966,0.019471066,0.017732201,0.032268047,0.83685184],"study_design_scores_gemma":[0.00012849919,0.00028922164,0.004811063,0.00031204452,0.0004377378,0.0005374627,0.00046413825,0.83058906,0.029049171,0.10728256,0.025995065,0.00010399396],"about_ca_topic_score_codex":0.0045253783,"about_ca_topic_score_gemma":0.011426297,"teacher_disagreement_score":0.005593807,"about_ca_system_score_codex":0.000890522,"about_ca_system_score_gemma":0.0017391507,"threshold_uncertainty_score":0.020230532},"labels":[],"label_agreement":null},{"id":"W3080121639","doi":"10.1016/j.metip.2020.100032","title":"Reducing the number of non-naïve participants in Mechanical Turk samples","year":2020,"lang":"en","type":"article","venue":"Methods in Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Psychology; Statistics; Mathematics","score_opus":0.2570277674644276,"score_gpt":0.5170510327490248,"score_spread":0.2600232652845971,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3080121639","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44191658,0.0024880953,0.36257178,0.008088436,0.003481058,0.12720351,0.00450668,0.003727734,0.046016086],"genre_scores_gemma":[0.55670744,0.00057534705,0.19790514,0.00949976,0.00090172054,0.22125202,0.0021773598,0.00065665855,0.01032465],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.91333085,0.059571356,0.0076690502,0.0060282676,0.010848682,0.002551867],"domain_scores_gemma":[0.82204777,0.103328325,0.010630948,0.03945228,0.021443933,0.003096654],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09318127,0.001548775,0.001729398,0.0017409539,0.004517208,0.003124069,0.0029680748,0.003210018,0.021641048],"category_scores_gemma":[0.22873242,0.0016520505,0.0011182581,0.0015707839,0.0042073433,0.003804295,0.0045862272,0.0020771963,0.010298138],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.020559825,0.012430485,0.16797487,0.009989399,0.0012444422,0.0030281662,0.07567944,0.0026690718,0.1019154,0.03621518,0.117182285,0.4511114],"study_design_scores_gemma":[0.00825065,0.016988669,0.36953583,0.0036139654,0.0013335937,0.0031242636,0.015815055,0.020919679,0.07757652,0.08778652,0.39422223,0.0008329982],"about_ca_topic_score_codex":0.0015477424,"about_ca_topic_score_gemma":0.005889078,"teacher_disagreement_score":0.90681875,"about_ca_system_score_codex":0.0007282454,"about_ca_system_score_gemma":0.0024555332,"threshold_uncertainty_score":0.49279553},"labels":[],"label_agreement":null},{"id":"W3080390953","doi":"10.24963/kr.2020/84","title":"Explainable and Argumentation-based Decision Making with Qualitative Preferences for Diagnostics and Prognostics of Alzheimer's Disease","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute on Aging; National Institute of Biomedical Imaging and Bioengineering; Canadian Institutes of Health Research; National Institutes of Health; Genentech; National Research Foundation Singapore; IXICO; H. Lundbeck A/S; Servier; Eisai; National Research Foundation; Northern California Institute for Research and Education; Pfizer; Biogen; BioClinica; F. Hoffmann-La Roche; University of Southern California; Eli Lilly and Company; U.S. Department of Defense; Meso Scale Diagnostics; Alzheimer's Disease Neuroimaging Initiative; Novartis Pharmaceuticals Corporation; Bristol-Myers Squibb; Foundation for the National Institutes of Health; Nanyang Technological University; Alzheimer's Association","keywords":"Argumentation theory; Computer science; Artificial intelligence; Machine learning; Semantic reasoner; Prognostics; Argumentative; Dialogical self; Management science; Data science; Data mining; Psychology; Epistemology; Engineering","score_opus":0.07094869352219021,"score_gpt":0.3348811080246179,"score_spread":0.26393241450242766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3080390953","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037470113,0.0007281738,0.95354545,0.0029275746,0.00006572341,0.0001894229,0.0003864489,0.0005204917,0.0041665775],"genre_scores_gemma":[0.53624403,0.0003831175,0.46115875,0.00024788574,0.00005762799,0.00021295788,0.00057793263,0.000052080624,0.0010656555],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9881413,0.008670833,0.0008310201,0.0008358566,0.0011926455,0.00032829875],"domain_scores_gemma":[0.9713955,0.02321586,0.0018690818,0.0012448452,0.001737698,0.0005369624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013416251,0.0010444814,0.00057451753,0.002791646,0.0013336227,0.003994258,0.0017154686,0.0022508383,0.0035595736],"category_scores_gemma":[0.03787244,0.00050223543,0.0019683545,0.0014432947,0.0019744337,0.00568965,0.0029051378,0.002652618,0.00044259717],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005965054,0.00032366216,0.009263283,0.0011077973,0.00033354876,0.00094642123,0.005047486,0.20706461,0.006201259,0.54365957,0.0048841243,0.22057186],"study_design_scores_gemma":[0.00008178501,0.00006931993,0.0014126499,0.00021615089,0.000118108284,0.00026546375,0.0007818995,0.58352655,0.0027604469,0.39851487,0.012175099,0.0000776573],"about_ca_topic_score_codex":0.0038473199,"about_ca_topic_score_gemma":0.0042528077,"teacher_disagreement_score":0.013416251,"about_ca_system_score_codex":0.0023207676,"about_ca_system_score_gemma":0.0022586188,"threshold_uncertainty_score":0.07095271},"labels":[],"label_agreement":null},{"id":"W3080528960","doi":"10.1037/xlm0000956","title":"Asymmetrical interference between item and order information in short-term memory.","year":2020,"lang":"en","type":"article","venue":"Journal of Experimental Psychology Learning Memory and Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; Natural Sciences and Engineering Research Council of Canada; National Institutes of Health","keywords":"Recall; Task (project management); Information retrieval; Computer science; Order (exchange); PsycINFO; Term (time); Serial position effect; Relation (database); Interference theory; Psychology; Free recall; Natural language processing; Cognitive psychology; Database; Working memory; Cognition","score_opus":0.04993582114645943,"score_gpt":0.32581343527448736,"score_spread":0.2758776141280279,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3080528960","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90137434,0.0013263951,0.07267373,0.0003525978,0.00013624772,0.00013068221,0.00021746302,0.0003588708,0.023429714],"genre_scores_gemma":[0.9832987,0.00030635018,0.013938235,0.00025021628,0.000053829775,0.00010017107,0.00033292393,0.000111377914,0.0016081373],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99756587,0.00063919474,0.00025842822,0.0004814851,0.00095185795,0.00010328936],"domain_scores_gemma":[0.97111064,0.0213477,0.002717223,0.0034550242,0.00076775416,0.00060162466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032927978,0.0005172271,0.00058867724,0.00070325396,0.00034551133,0.00153453,0.0010979259,0.0006797247,0.0036890358],"category_scores_gemma":[0.030751277,0.0004811493,0.00034813018,0.00064508093,0.0012314316,0.003803327,0.0021847144,0.0008673382,0.0008314256],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005214689,0.0007087487,0.056254257,0.0017158615,0.0003233921,0.00080392975,0.005741127,0.003150194,0.51708555,0.043406174,0.0025772557,0.36301884],"study_design_scores_gemma":[0.0009869329,0.00422358,0.45413598,0.00052903546,0.0008596929,0.005219745,0.0017206104,0.06312588,0.2142167,0.241038,0.013663102,0.00028073505],"about_ca_topic_score_codex":0.0005565343,"about_ca_topic_score_gemma":0.00087958615,"teacher_disagreement_score":0.0036890358,"about_ca_system_score_codex":0.0004940192,"about_ca_system_score_gemma":0.00046448462,"threshold_uncertainty_score":0.017414212},"labels":[],"label_agreement":null},{"id":"W3080788362","doi":"10.21742/ijseia.2020.14.1.02","title":"Chatbot Analytics Based on Question Answering System and Deep Learning: Case Study for Movie Smart Automatic Answering","year":2020,"lang":"en","type":"article","venue":"International Journal of Software Engineering and Its Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Lakehead University","keywords":"Chatbot; Computer science; Question answering; Artificial intelligence; Natural language processing; Deep learning; Word embedding; Recurrent neural network; Sentiment analysis; Natural language; Natural language understanding; Artificial neural network; Embedding","score_opus":0.01723132915862683,"score_gpt":0.2574833423538425,"score_spread":0.2402520131952157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3080788362","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8750381,0.0020297323,0.08968967,0.0021733546,0.0002096594,0.0012300082,0.0020441443,0.012442836,0.015142438],"genre_scores_gemma":[0.934305,0.00029067075,0.0546827,0.00039614606,0.00006986171,0.0002450781,0.0017205691,0.00021347265,0.008076405],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99775726,0.001095183,0.000113197275,0.00033949292,0.00044630648,0.0002485987],"domain_scores_gemma":[0.99488896,0.0030358618,0.00017510017,0.000525495,0.0008334226,0.0005412416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022065288,0.00072110945,0.0007295315,0.0008312573,0.00093372696,0.0011394738,0.0014930448,0.0017459057,0.0030678557],"category_scores_gemma":[0.0058874255,0.00025664966,0.0005151405,0.00088993885,0.00070099096,0.0022026028,0.0012487366,0.0012753136,0.0013111413],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0066811056,0.008090694,0.067675434,0.0042263498,0.00049889716,0.025004104,0.017029453,0.05295544,0.08712608,0.02253015,0.08458486,0.6235975],"study_design_scores_gemma":[0.00041250658,0.002581788,0.037984625,0.00020890107,0.00019592182,0.0038932692,0.007327204,0.8023941,0.06451821,0.0095498655,0.070766516,0.00016710251],"about_ca_topic_score_codex":0.009210084,"about_ca_topic_score_gemma":0.011480312,"teacher_disagreement_score":0.009210084,"about_ca_system_score_codex":0.0012225792,"about_ca_system_score_gemma":0.00092274137,"threshold_uncertainty_score":0.018312931},"labels":[],"label_agreement":null},{"id":"W3084790339","doi":"10.1007/978-3-030-58219-7_7","title":"Argument Retrieval from Web","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Canadian Institute of Steel Construction","keywords":"Argumentative; Computer science; Information retrieval; PageRank; Search engine; Relevance (law); Argument (complex analysis); Task (project management); World Wide Web","score_opus":0.023068733615182305,"score_gpt":0.2370536448138783,"score_spread":0.213984911198696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3084790339","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.097943306,0.016331188,0.6221437,0.004416241,0.0013124419,0.0007130799,0.017045824,0.027692191,0.21240205],"genre_scores_gemma":[0.45393637,0.0076969434,0.3523134,0.0009484292,0.0010347201,0.00041509935,0.04964173,0.0029968726,0.13101642],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999255,0.00021968444,0.000047027745,0.000121871046,0.0002970964,0.000059233254],"domain_scores_gemma":[0.99867,0.0008262376,0.00005069949,0.00023153918,0.0001808597,0.000040656138],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006248974,0.00077081635,0.0009644286,0.0038583938,0.000675072,0.0028397175,0.000903892,0.0011992296,0.03464388],"category_scores_gemma":[0.0058756126,0.00040074743,0.00081288157,0.0029737258,0.00037781455,0.004886667,0.0017629476,0.0011830723,0.021641305],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045481627,0.00027276165,0.0007465722,0.0008686462,0.00010562747,0.00041317334,0.0003764674,0.0054669017,0.015856069,0.037507463,0.13729124,0.8006402],"study_design_scores_gemma":[0.00018341297,0.00026485874,0.0034614783,0.00053754286,0.00031554425,0.0015377321,0.0009872469,0.31246865,0.046768766,0.22913639,0.40422288,0.000115428],"about_ca_topic_score_codex":0.0008060153,"about_ca_topic_score_gemma":0.0016005253,"teacher_disagreement_score":0.03464388,"about_ca_system_score_codex":0.00062045985,"about_ca_system_score_gemma":0.0005784836,"threshold_uncertainty_score":0.11589533},"labels":[],"label_agreement":null},{"id":"W3086286810","doi":"10.1007/978-3-031-02178-7_5","title":"Semantic Relations and Deep Learning","year":2021,"lang":"en","type":"book-chapter","venue":"Synthesis lectures on human language technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Permission; Relation (database); Computer science; Artificial intelligence; Deep learning; Natural language processing; Cognitive science; Psychology; Philosophy; Epistemology; Data mining","score_opus":0.020797629303994487,"score_gpt":0.2525003954038954,"score_spread":0.23170276609990093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3086286810","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010233533,0.032999575,0.8077251,0.006916762,0.001629908,0.000033213135,0.00093268685,0.0016560098,0.13787326],"genre_scores_gemma":[0.29453415,0.036052186,0.30303013,0.0013594241,0.0026173564,0.00018094108,0.0040884283,0.0013501333,0.35678723],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9998061,0.000040007777,0.000010253596,0.000056731114,0.000068339476,0.000018648441],"domain_scores_gemma":[0.9996451,0.00021672065,0.000017901462,0.00005701941,0.000045342716,0.00001799022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004006737,0.0007224649,0.00053796434,0.0007700774,0.00036438904,0.0017908555,0.0006313776,0.0006608895,0.012153878],"category_scores_gemma":[0.0014237311,0.00050081685,0.00040521592,0.0015734612,0.0010288919,0.0042634928,0.0010102595,0.0021444755,0.0032895193],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029750237,0.000033597946,0.00017221295,0.00021015404,0.000028779314,0.000025563577,0.000112837624,0.013577601,0.0015611002,0.57695514,0.05245186,0.3548414],"study_design_scores_gemma":[0.0000028124846,0.0000067964743,0.0001902155,0.000047213973,0.000010985719,0.00003483372,0.000023693945,0.032103367,0.0010873306,0.90609133,0.06039295,0.000008565336],"about_ca_topic_score_codex":0.0017806514,"about_ca_topic_score_gemma":0.0030073712,"teacher_disagreement_score":0.012153878,"about_ca_system_score_codex":0.0009508686,"about_ca_system_score_gemma":0.0005309527,"threshold_uncertainty_score":0.040658712},"labels":[],"label_agreement":null},{"id":"W3086996502","doi":"10.18653/v1/2021.findings-acl.378","title":"MLMLM: Link Prediction with Mean Likelihood Masked Language Model","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute; Canadian Institute for Advanced Research","funders":"Canada First Research Excellence Fund; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Interpretability; Computer science; Scalability; Language model; Artificial intelligence; Link (geometry); Verifiable secret sharing; Embedding; Machine learning; Scale (ratio); Data mining; Natural language processing; Database","score_opus":0.01731573470158831,"score_gpt":0.23185099937216969,"score_spread":0.21453526467058137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3086996502","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025004817,0.0025983616,0.9168354,0.0018346518,0.00035668453,0.00035138012,0.013211533,0.036535237,0.0032719104],"genre_scores_gemma":[0.30980444,0.0010734188,0.62877774,0.0012540296,0.00052856613,0.0008258228,0.049054638,0.0019027736,0.006778634],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971559,0.0011278901,0.00017989708,0.0008065862,0.0005681473,0.00016155485],"domain_scores_gemma":[0.9927355,0.004960173,0.0003329548,0.0012937878,0.0005040584,0.0001736224],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035692877,0.0019761813,0.0014183066,0.0029054459,0.00090449327,0.0027246561,0.0038720237,0.0024395555,0.006376645],"category_scores_gemma":[0.016619163,0.00092776713,0.0022634647,0.002701446,0.00079017685,0.0051021078,0.003196124,0.0039189984,0.0061257007],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011850725,0.0006092838,0.009058519,0.0012572851,0.0007620367,0.0007863392,0.0005968007,0.31894118,0.0070492574,0.030876046,0.12647654,0.50240165],"study_design_scores_gemma":[0.000078367644,0.00006886455,0.00043861833,0.000042177468,0.00004789303,0.00010396889,0.000048155274,0.9489399,0.00199868,0.041991152,0.0062087784,0.000033356948],"about_ca_topic_score_codex":0.009492472,"about_ca_topic_score_gemma":0.0139056025,"teacher_disagreement_score":0.009492472,"about_ca_system_score_codex":0.0012752047,"about_ca_system_score_gemma":0.002216531,"threshold_uncertainty_score":0.021332026},"labels":[],"label_agreement":null},{"id":"W3087421864","doi":"10.1007/s40593-020-00211-5","title":"Automated Essay Scoring and the Deep Learning Black Box: How Are Rubric Scores Determined?","year":2020,"lang":"en","type":"article","venue":"International Journal of Artificial Intelligence in Education","topic":"Topic Modeling","field":"Computer Science","cited_by":85,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Rubric; Grading (engineering); CONTEST; Computer science; Artificial intelligence; Deep learning; Machine learning; Natural language processing; Mathematics education; Psychology","score_opus":0.038663089128390074,"score_gpt":0.3147992035990448,"score_spread":0.2761361144706547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3087421864","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36000264,0.004416818,0.55281067,0.007816601,0.001183056,0.00055367086,0.005856837,0.014156181,0.05320356],"genre_scores_gemma":[0.80181664,0.00057957787,0.18093963,0.00067865127,0.00027751358,0.00038378846,0.0032353525,0.00075878453,0.011330148],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9818703,0.009970119,0.0011639004,0.0017752644,0.004545251,0.0006752276],"domain_scores_gemma":[0.9302749,0.033256203,0.005223432,0.0073765083,0.022415144,0.0014538128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014809271,0.000900137,0.0014471256,0.0033945453,0.00072887965,0.0047222,0.001834864,0.0012688066,0.007214185],"category_scores_gemma":[0.09970408,0.0005063989,0.00041454684,0.002534917,0.0006984776,0.0034808384,0.0020367478,0.0019435569,0.0076452186],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004058827,0.00034636245,0.05677023,0.00022176045,0.00015317494,0.000026845519,0.00032211968,0.004726457,0.0033306207,0.0040901457,0.029089011,0.90051734],"study_design_scores_gemma":[0.00027593254,0.0008767145,0.18003264,0.001187358,0.00025787405,0.0003946137,0.0018776838,0.66040355,0.03864,0.06368503,0.051987275,0.0003813698],"about_ca_topic_score_codex":0.0043960265,"about_ca_topic_score_gemma":0.00823978,"teacher_disagreement_score":0.014809271,"about_ca_system_score_codex":0.0009939474,"about_ca_system_score_gemma":0.002284458,"threshold_uncertainty_score":0.07831985},"labels":[],"label_agreement":null},{"id":"W3088218728","doi":"10.1017/9781108922036","title":"Can We Be Wrong? The Problem of Textual Evidence in a Time of Data","year":2020,"lang":"en","type":"book","venue":"Cambridge University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Generalizability theory; Generalization; Element (criminal law); Computer science; Field (mathematics); Data science; Artificial intelligence; Discipline; Epistemology; Natural language processing; Psychology; Sociology; Mathematics; Social science; Political science; Philosophy","score_opus":0.08577631419635601,"score_gpt":0.2379681552953743,"score_spread":0.1521918410990183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3088218728","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009372923,0.034987945,0.48850143,0.34366056,0.0052306885,0.00020662624,0.0015626176,0.0004668005,0.11601047],"genre_scores_gemma":[0.4389028,0.03134421,0.37490386,0.060692336,0.016261568,0.0011829743,0.0016528054,0.0015520452,0.07350742],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9787629,0.013887708,0.00084199995,0.002187829,0.004030026,0.00028954737],"domain_scores_gemma":[0.829283,0.15613832,0.0033134057,0.0071931323,0.0033625618,0.00070955994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030416155,0.0008705815,0.0018907702,0.005184226,0.0036143747,0.015597612,0.0036889727,0.0055376035,0.008477568],"category_scores_gemma":[0.12328605,0.0011724958,0.0012015324,0.005641414,0.025355726,0.047016114,0.004650903,0.011055636,0.0022218663],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003618309,0.0000144553105,0.0006198967,0.00038003924,0.00007658448,0.0001507472,0.003243818,0.0011139445,0.00009875981,0.91540575,0.019856052,0.059003893],"study_design_scores_gemma":[0.0000059654235,0.0000094898605,0.00017012031,0.00040517806,0.000015988955,0.00015633256,0.0010951755,0.0022504313,0.00012813671,0.94854975,0.04719292,0.00002044911],"about_ca_topic_score_codex":0.0021057043,"about_ca_topic_score_gemma":0.002630534,"teacher_disagreement_score":0.030416155,"about_ca_system_score_codex":0.003726602,"about_ca_system_score_gemma":0.0035601943,"threshold_uncertainty_score":0.16085798},"labels":[],"label_agreement":null},{"id":"W3089160253","doi":"","title":"Semantic Similarity Frontiers: From Concepts to Documents","year":2015,"lang":"en","type":"article","venue":"Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Semantic similarity; Similarity (geometry); Semantics (computer science); Natural language processing; Information retrieval; Sentence; SemEval; Artificial intelligence; Word (group theory); Semantic computing; Component (thermodynamics); Semantic Web; Linguistics","score_opus":0.07529857966887143,"score_gpt":0.46713590103802416,"score_spread":0.3918373213691527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089160253","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018297687,0.17468423,0.6704322,0.03950157,0.003943226,0.00082345377,0.0051009576,0.0025080831,0.08470862],"genre_scores_gemma":[0.26131946,0.075438574,0.6234647,0.0059491494,0.0053607333,0.0017027092,0.007486814,0.0013465519,0.017931258],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98896503,0.0048953597,0.0009885584,0.001958112,0.0027789874,0.00041399366],"domain_scores_gemma":[0.9829217,0.012517802,0.00066836114,0.0015278409,0.0017515738,0.0006126796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007030482,0.0016750013,0.0019059106,0.01577358,0.0022697148,0.017457986,0.003467682,0.0032737602,0.016323077],"category_scores_gemma":[0.041531414,0.0010694809,0.0018536847,0.015969729,0.009744746,0.044080645,0.008495122,0.0047148303,0.004286126],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012837608,0.000062677784,0.0009319545,0.0020403434,0.00009078026,0.00021759259,0.0043061464,0.0020564303,0.00083508954,0.70620435,0.02874314,0.25438312],"study_design_scores_gemma":[0.000019068497,0.00003541071,0.000801977,0.0006318014,0.00002357598,0.00035815188,0.0019291028,0.006017998,0.0005252186,0.8562193,0.13340206,0.00003633168],"about_ca_topic_score_codex":0.0040528555,"about_ca_topic_score_gemma":0.0018162557,"teacher_disagreement_score":0.017457986,"about_ca_system_score_codex":0.004879111,"about_ca_system_score_gemma":0.003480307,"threshold_uncertainty_score":0.0546062},"labels":[],"label_agreement":null},{"id":"W3089233166","doi":"10.18653/v1/2020.emnlp-main.97","title":"SSMBA: Self-Supervised Manifold Based Data Augmentation for Improving Out-of-Domain Robustness","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Samsung; Vector Institute; Government of Canada; Samsung Advanced Institute of Technology; Canadian Institute for Advanced Research","keywords":"Overfitting; Robustness (evolution); Computer science; Artificial intelligence; Domain (mathematical analysis); Training set; Generalization; Machine learning; Natural language; Manifold (fluid mechanics); Labeled data; Natural language processing; Mathematics; Artificial neural network","score_opus":0.10664366249065856,"score_gpt":0.30480282316480883,"score_spread":0.1981591606741503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089233166","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045211818,0.0016501884,0.9234469,0.0006184324,0.0003204294,0.0002324602,0.0016651497,0.024878634,0.0019759806],"genre_scores_gemma":[0.4064551,0.0007113385,0.56936264,0.0008071297,0.0002988778,0.0006927972,0.013824066,0.0021971879,0.0056508128],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971982,0.0011728051,0.00018911873,0.00077120314,0.000521744,0.00014692893],"domain_scores_gemma":[0.99330235,0.002548504,0.00052888197,0.0024088158,0.0009654584,0.00024615994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003512463,0.0025559706,0.0017913876,0.0020457727,0.0008301491,0.0015516861,0.0031179918,0.002068539,0.002974611],"category_scores_gemma":[0.013704951,0.0008223174,0.002208003,0.0015136524,0.0016741117,0.0032833559,0.0040395544,0.0035751623,0.003582571],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007996428,0.00068250205,0.0063896277,0.000849017,0.0005553512,0.00036138995,0.0007052759,0.23740664,0.042329982,0.0073827044,0.050109595,0.6524283],"study_design_scores_gemma":[0.00004565741,0.00018522845,0.0008446513,0.000045250716,0.000037527258,0.00016140897,0.000068895606,0.9709542,0.01445381,0.0071989084,0.0059636496,0.00004081961],"about_ca_topic_score_codex":0.002845145,"about_ca_topic_score_gemma":0.0043956656,"teacher_disagreement_score":0.003512463,"about_ca_system_score_codex":0.00081281964,"about_ca_system_score_gemma":0.001122878,"threshold_uncertainty_score":0.018575847},"labels":[],"label_agreement":null},{"id":"W3089306998","doi":"10.48550/arxiv.2009.12452","title":"BET: A Backtranslation Approach for Easy Data Augmentation in Transformer-based Paraphrase Identification Context","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Paraphrase; Generalizability theory; Computer science; Transformer; Artificial intelligence; Natural language processing; Deep learning; Training set; Machine learning; Statistics; Mathematics","score_opus":0.2658474981201677,"score_gpt":0.24292606176648887,"score_spread":0.022921436353678853,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089306998","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15583827,0.0025531687,0.7576544,0.0013256955,0.0006674359,0.0007103751,0.013497274,0.058075003,0.009678421],"genre_scores_gemma":[0.47009304,0.0006359522,0.48101884,0.00075620145,0.00026243654,0.0010266305,0.037750423,0.0017787175,0.0066777077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99848926,0.00056584494,0.00017371048,0.00040824304,0.00024354126,0.00011939345],"domain_scores_gemma":[0.9963337,0.0011975648,0.00018325726,0.0014872398,0.00067560596,0.00012266781],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026290833,0.0017130516,0.0009283327,0.0024673925,0.00080695096,0.0016253213,0.0026282382,0.0015127043,0.007574538],"category_scores_gemma":[0.012632012,0.000636996,0.0012071674,0.0022753251,0.0007138712,0.0046107364,0.0037930773,0.0027938331,0.0060350266],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010643874,0.00057176664,0.005208217,0.0007738585,0.00025960448,0.00041765012,0.0007339466,0.036046896,0.03344732,0.008258345,0.04385154,0.8693664],"study_design_scores_gemma":[0.00023744973,0.0006432495,0.004195618,0.00013675557,0.00015231343,0.0007030452,0.00052468537,0.878075,0.050292347,0.024601975,0.040330116,0.000107408465],"about_ca_topic_score_codex":0.0029363253,"about_ca_topic_score_gemma":0.0074699246,"teacher_disagreement_score":0.007574538,"about_ca_system_score_codex":0.0006526918,"about_ca_system_score_gemma":0.0012805937,"threshold_uncertainty_score":0.025339305},"labels":[],"label_agreement":null},{"id":"W3089901159","doi":"10.1145/3377812.3390790","title":"Semantic analysis of issues on Google play and Twitter","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; Kelowna General Hospital; University of British Columbia","funders":"Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; World Wide Web; App store; Semantics (computer science); Mobile apps; Information retrieval; Social media; Resource (disambiguation); Sentiment analysis; Semantic analysis (machine learning); Data science; Natural language processing","score_opus":0.04085643520667663,"score_gpt":0.2763065765115711,"score_spread":0.23545014130489444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089901159","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9272148,0.0018404734,0.02629236,0.0015942439,0.0003153779,0.00058962614,0.022379376,0.0011850381,0.018588657],"genre_scores_gemma":[0.9618568,0.0006937061,0.015557963,0.00011973524,0.00029781798,0.00041651292,0.017100316,0.00012339876,0.003833834],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9988715,0.00026881296,0.00011887507,0.00014474566,0.0004904988,0.00010562909],"domain_scores_gemma":[0.9964173,0.0017483737,0.0006416637,0.00012102866,0.0009393018,0.00013229091],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006476899,0.00058685377,0.00031238998,0.0085646,0.0009180973,0.0012857352,0.0003045137,0.00047688987,0.0011081742],"category_scores_gemma":[0.0051283496,0.0001357554,0.0005892026,0.0048929756,0.00034923712,0.0021730897,0.00076075905,0.00040217183,0.0006725172],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021427122,0.00045464723,0.31735635,0.0041276896,0.00043957986,0.0054490073,0.035968855,0.008701097,0.095673434,0.02056984,0.07185069,0.43726614],"study_design_scores_gemma":[0.000048054586,0.0002960728,0.6609682,0.00036938934,0.00029131369,0.0018135007,0.031231156,0.1405591,0.022729063,0.009406317,0.1320906,0.00019711081],"about_ca_topic_score_codex":0.008757392,"about_ca_topic_score_gemma":0.013975009,"teacher_disagreement_score":0.008757392,"about_ca_system_score_codex":0.0009378648,"about_ca_system_score_gemma":0.0005930648,"threshold_uncertainty_score":0.017412841},"labels":[],"label_agreement":null},{"id":"W3089922444","doi":"10.3233/web-200445","title":"Document classification using convolutional neural networks with small window sizes and latent semantic analysis","year":2020,"lang":"en","type":"article","venue":"Web Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Convolutional neural network; Word2vec; Artificial intelligence; Leverage (statistics); Latent semantic analysis; Pattern recognition (psychology); Sliding window protocol; Window (computing); Word (group theory); Robustness (evolution); Document classification; Mathematics","score_opus":0.07204337661805335,"score_gpt":0.2603511661477368,"score_spread":0.18830778952968344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089922444","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17948267,0.002157126,0.80169827,0.00046247157,0.0003440988,0.00020579073,0.0021101024,0.008422759,0.0051167556],"genre_scores_gemma":[0.64085096,0.0009357129,0.34368268,0.00016636236,0.00016141038,0.0001743434,0.0051308176,0.00019461027,0.008703009],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995838,0.000048317623,0.00003801793,0.00013852654,0.00011791122,0.00007334818],"domain_scores_gemma":[0.9994363,0.00018307494,0.00008050701,0.00010540436,0.00015980343,0.00003491648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007532337,0.0010745105,0.000540643,0.0015360132,0.00031302977,0.0011148215,0.0007696779,0.000588228,0.0018418622],"category_scores_gemma":[0.0016969555,0.00029502602,0.0007694282,0.0015462603,0.0002773072,0.0019911884,0.0006818443,0.0010652685,0.0013661564],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004550824,0.00030180873,0.005028702,0.00022333379,0.00021995856,0.00014907659,0.00010032812,0.09798902,0.045987487,0.0050904853,0.010268877,0.83418584],"study_design_scores_gemma":[0.000012356516,0.00006108469,0.0014263495,0.000015093705,0.000029387853,0.000035371017,0.000020541458,0.9823167,0.011291761,0.0028216778,0.001954932,0.000014737818],"about_ca_topic_score_codex":0.011920165,"about_ca_topic_score_gemma":0.016527355,"teacher_disagreement_score":0.011920165,"about_ca_system_score_codex":0.0012585135,"about_ca_system_score_gemma":0.0008259715,"threshold_uncertainty_score":0.023701549},"labels":[],"label_agreement":null},{"id":"W3090223312","doi":"10.5539/cis.v13n4p1","title":"Arabic-to-Malay Machine Translation Using Transfer Approach","year":2020,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Malay; Computer science; Machine translation; Natural language processing; Arabic; Artificial intelligence; BLEU; Translation (biology); n-gram; Example-based machine translation; Meaning (existential); Machine translation software usability; Rule-based machine translation; Linguistics; Language model","score_opus":0.05363112321334172,"score_gpt":0.25133309024215483,"score_spread":0.1977019670288131,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3090223312","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07243877,0.0012724728,0.89162207,0.0006418429,0.000300412,0.00044447242,0.0007499003,0.011081392,0.021448635],"genre_scores_gemma":[0.66126716,0.0012522361,0.31216562,0.00021237854,0.00019746042,0.00039805792,0.0018659775,0.0004288336,0.022212276],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995028,0.00018558992,0.000044886,0.000118614196,0.0001066816,0.000041401905],"domain_scores_gemma":[0.99957365,0.00015933809,0.000029464549,0.00009373687,0.00012557935,0.000018286373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00059170666,0.0005959578,0.00047903095,0.00090085604,0.00062550197,0.00078242633,0.00046529042,0.000498366,0.0058269496],"category_scores_gemma":[0.0013011675,0.00015529688,0.00055268774,0.0009344821,0.00026534588,0.0013597478,0.0008665939,0.0005622135,0.0044551254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023730782,0.00021257739,0.0013987498,0.00047243046,0.00010416865,0.00057257497,0.00055086135,0.023028739,0.05555954,0.01078168,0.008582525,0.89849883],"study_design_scores_gemma":[0.00006867901,0.00044401595,0.0046445066,0.00007692843,0.00015822292,0.0013335649,0.00067388767,0.7777211,0.12675819,0.03388291,0.054137178,0.000100832876],"about_ca_topic_score_codex":0.0015785386,"about_ca_topic_score_gemma":0.0016123764,"teacher_disagreement_score":0.0058269496,"about_ca_system_score_codex":0.0003719642,"about_ca_system_score_gemma":0.00081392674,"threshold_uncertainty_score":0.019493103},"labels":[],"label_agreement":null},{"id":"W3091315598","doi":"10.1109/ijcnn48605.2020.9207289","title":"Transformer Decoder Based Reinforcement Learning Approach for Conversational Response Generation","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Transformer; Conversation; Reinforcement learning; Artificial intelligence; Language model; Speech recognition; Natural language processing; Machine learning; Voltage; Engineering","score_opus":0.06948471122514176,"score_gpt":0.25953214228844246,"score_spread":0.1900474310633007,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091315598","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018012468,0.00023378902,0.9773142,0.00031398723,0.00005883004,0.00010284681,0.0000851231,0.0015389618,0.0023397917],"genre_scores_gemma":[0.8152251,0.0002002874,0.17481536,0.00036026986,0.00007736863,0.00035847054,0.00029582594,0.00024526386,0.008422026],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99873155,0.0006829074,0.000048924983,0.00027387313,0.00017320436,0.00008946886],"domain_scores_gemma":[0.9975637,0.0016735818,0.00011194151,0.00016790748,0.00033966437,0.00014326471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021360961,0.0007836956,0.00085941685,0.0005286109,0.00038114816,0.0006907511,0.001819265,0.0012039429,0.0039800066],"category_scores_gemma":[0.0062741735,0.00039844925,0.0006254325,0.00035741777,0.0008281635,0.001378591,0.0013253638,0.0019116956,0.0013397653],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005813607,0.00041137676,0.001786111,0.00024014372,0.00010132256,0.00027315487,0.0007025568,0.67336464,0.017732998,0.029319381,0.0046293903,0.2708576],"study_design_scores_gemma":[0.000012576337,0.000047715257,0.00005212384,0.0000035579776,0.000007296863,0.00002339106,0.000012045003,0.99351126,0.0013973129,0.004377216,0.0005486692,0.000006845015],"about_ca_topic_score_codex":0.0032002144,"about_ca_topic_score_gemma":0.0035697222,"teacher_disagreement_score":0.0039800066,"about_ca_system_score_codex":0.0009748608,"about_ca_system_score_gemma":0.0011917198,"threshold_uncertainty_score":0.013314486},"labels":[],"label_agreement":null},{"id":"W3091829090","doi":"10.48550/arxiv.2010.02838","title":"A Closer Look at Codistillation for Distributed Training","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Training (meteorology); Computer science; Geography","score_opus":0.17539130866281383,"score_gpt":0.2065511311427058,"score_spread":0.031159822479891963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091829090","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033217117,0.0013609185,0.94959295,0.0045765853,0.00035147957,0.00011272177,0.00015693091,0.0039423746,0.006688966],"genre_scores_gemma":[0.6033665,0.0007759271,0.38375503,0.0017263326,0.00045108824,0.000383291,0.00063652656,0.0014841545,0.007421176],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99663,0.0014854586,0.00013657866,0.0009933902,0.00046038273,0.00029416793],"domain_scores_gemma":[0.9894288,0.004522581,0.0004426942,0.0044111595,0.00075126643,0.00044357748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063744523,0.0011426776,0.0015027878,0.0004732454,0.0011270108,0.0029590295,0.00370131,0.0019930596,0.0050574206],"category_scores_gemma":[0.01980305,0.0006413516,0.00096387835,0.00088736817,0.0019603344,0.0059361937,0.003857109,0.006177212,0.0013872511],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010929925,0.0006010822,0.006857711,0.00045970725,0.00033399634,0.0004601999,0.0012240789,0.52397966,0.023812769,0.15929823,0.02789399,0.25398564],"study_design_scores_gemma":[0.00010160725,0.00012211535,0.0006509966,0.000044471588,0.000022636103,0.000105516876,0.000101305086,0.9264311,0.005110128,0.057691783,0.009587313,0.000031007054],"about_ca_topic_score_codex":0.0059354203,"about_ca_topic_score_gemma":0.009844942,"teacher_disagreement_score":0.0063744523,"about_ca_system_score_codex":0.0019578275,"about_ca_system_score_gemma":0.0023029626,"threshold_uncertainty_score":0.03371173},"labels":[],"label_agreement":null},{"id":"W3091895957","doi":"10.1007/978-3-030-60887-3_19","title":"Evaluation of Similarity Measures in a Benchmark for Spanish Paraphrasing Detection","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Paraphrase; Similarity (geometry); Computer science; Natural language processing; Artificial intelligence; Vocabulary; Benchmark (surveying); Semantic similarity; Boundary (topology); Computation; Information retrieval; Linguistics; Algorithm; Mathematics","score_opus":0.0693253929892796,"score_gpt":0.29153093333462454,"score_spread":0.22220554034534495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091895957","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.825232,0.010939461,0.086589016,0.00060827786,0.0012838303,0.0018244176,0.021499466,0.028413609,0.023609836],"genre_scores_gemma":[0.70531315,0.0017120972,0.17984253,0.00031149414,0.0003863945,0.000893425,0.10011677,0.0018081248,0.009616056],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9933775,0.0020224017,0.0009971136,0.0015727386,0.0016400397,0.000390259],"domain_scores_gemma":[0.98272055,0.007632533,0.00077803177,0.0023512996,0.0052598375,0.0012576631],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048957323,0.0023121356,0.0018498474,0.007838544,0.0014771065,0.003183865,0.0027707876,0.0028376724,0.0067309546],"category_scores_gemma":[0.021740275,0.0003636499,0.0010751336,0.005231481,0.0006317583,0.0034717547,0.0029999444,0.0012199072,0.005757358],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005604555,0.0042245993,0.016937781,0.004581843,0.0011084629,0.0010935114,0.0012446297,0.016529897,0.060133964,0.0028073143,0.059237055,0.8264963],"study_design_scores_gemma":[0.0026215645,0.008548601,0.09561217,0.00063072005,0.0011556955,0.004464874,0.0051516746,0.6728785,0.14494722,0.007606424,0.05594873,0.00043390732],"about_ca_topic_score_codex":0.0067182654,"about_ca_topic_score_gemma":0.008006216,"teacher_disagreement_score":0.007838544,"about_ca_system_score_codex":0.0010077428,"about_ca_system_score_gemma":0.0012548269,"threshold_uncertainty_score":0.025891423},"labels":[],"label_agreement":null},{"id":"W3091948559","doi":"10.1007/978-3-030-59830-3_29","title":"Predicting US Elections with Social Media and Neural Networks","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Artificial neural network; Social media; Artificial intelligence; World Wide Web","score_opus":0.01842359035266617,"score_gpt":0.22315437148536824,"score_spread":0.20473078113270207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3091948559","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7695846,0.009771404,0.18059023,0.0036178872,0.0010450933,0.00016668346,0.0063789748,0.0017131822,0.027131995],"genre_scores_gemma":[0.9671765,0.001126828,0.019325294,0.00012030987,0.0009030946,0.000055305743,0.002819645,0.00007209983,0.008400895],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997148,0.00010965652,0.000015208852,0.000069233414,0.00004131156,0.000049659615],"domain_scores_gemma":[0.99802196,0.0015971585,0.00012337846,0.000077472905,0.00011728479,0.00006275575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00096388435,0.00077081344,0.00064566586,0.001958848,0.00035933458,0.0013002892,0.0005780556,0.00076794927,0.004181398],"category_scores_gemma":[0.004537906,0.0004776808,0.000775938,0.0015262135,0.0002433036,0.0021107302,0.0004825318,0.0012474344,0.0015553982],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008216521,0.00066783,0.07325427,0.000220748,0.0005474798,0.00013626693,0.00009795675,0.44421023,0.0014189945,0.008993683,0.033878546,0.4357523],"study_design_scores_gemma":[0.000009746148,0.000021600494,0.005166169,0.000012753398,0.000026835423,0.000015019366,0.000023857543,0.9868951,0.00026826482,0.00665075,0.00090350525,0.000006387554],"about_ca_topic_score_codex":0.00975568,"about_ca_topic_score_gemma":0.023413518,"teacher_disagreement_score":0.00975568,"about_ca_system_score_codex":0.0007531978,"about_ca_system_score_gemma":0.00029278465,"threshold_uncertainty_score":0.019397795},"labels":[],"label_agreement":null},{"id":"W3092009239","doi":"10.1007/s11192-020-03718-9","title":"Navigation-based candidate expansion and pretrained language models for citation recommendation","year":2020,"lang":"en","type":"article","venue":"Scientometrics","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; University of Waterloo","funders":"","keywords":"Computer science; Benchmark (surveying); Citation; Ranking (information retrieval); Vocabulary; Domain (mathematical analysis); Task (project management); Information retrieval; Artificial intelligence; Language model; Machine learning; Natural language processing; World Wide Web; Linguistics","score_opus":0.06364737314638658,"score_gpt":0.31168752864163235,"score_spread":0.24804015549524577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092009239","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11115035,0.0058788774,0.86421907,0.0026196633,0.00059926754,0.00025698508,0.003628862,0.007662918,0.003983981],"genre_scores_gemma":[0.73483866,0.0026766455,0.23262276,0.00084557297,0.0014558915,0.00069436664,0.011815393,0.00068537967,0.014365306],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99808085,0.0007723055,0.00018914719,0.00043410665,0.00030446606,0.00021912914],"domain_scores_gemma":[0.9907673,0.0068277344,0.000337206,0.00061879057,0.0011594144,0.000289421],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0035690896,0.0015197905,0.002557183,0.0058532176,0.0011926129,0.0021473067,0.0029142268,0.0029115041,0.0045795557],"category_scores_gemma":[0.016086888,0.00075264275,0.0023656378,0.006114601,0.00064951123,0.00400384,0.0017666767,0.0030284626,0.004379843],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012859902,0.0011338298,0.008003539,0.0005970355,0.00054447266,0.0003927727,0.00034628532,0.26799837,0.0063639614,0.014617704,0.033464987,0.665251],"study_design_scores_gemma":[0.000027534745,0.000034459918,0.00037054365,0.000015934096,0.000053878135,0.000041954132,0.000019684107,0.99274534,0.0006637817,0.0052971332,0.000711868,0.0000179703],"about_ca_topic_score_codex":0.021008493,"about_ca_topic_score_gemma":0.031167516,"teacher_disagreement_score":0.99414676,"about_ca_system_score_codex":0.0011835418,"about_ca_system_score_gemma":0.0028121676,"threshold_uncertainty_score":0.041772425},"labels":[],"label_agreement":null},{"id":"W3092108395","doi":"","title":"Using Terminology and a Concept Hierarchy for Restricted-Domain Question-Answering","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Terminology; Hierarchy; Computer science; Domain (mathematical analysis); Question answering; Natural language processing; Artificial intelligence; Linguistics; Mathematics; Philosophy; Political science","score_opus":0.03269877877566846,"score_gpt":0.28827980105954654,"score_spread":0.2555810222838781,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092108395","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015162528,0.0009302766,0.9809553,0.0004444931,0.00003919648,0.00011290368,0.00018383579,0.00065908633,0.0015123774],"genre_scores_gemma":[0.22742419,0.00048912235,0.77006555,0.00020474984,0.000105970576,0.00024148102,0.00072199764,0.000116354386,0.00063060416],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9918607,0.005694152,0.00040582905,0.0008523664,0.0010027721,0.00018424453],"domain_scores_gemma":[0.9817328,0.012670878,0.0009910582,0.0029585527,0.0014057368,0.00024104567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008384698,0.00072262663,0.00097621174,0.0045914175,0.0009722368,0.002391278,0.0014539846,0.0014314122,0.0019959745],"category_scores_gemma":[0.037574027,0.0007577376,0.0015332177,0.0036082459,0.0019108716,0.0070053106,0.0027886452,0.0017120382,0.0009434237],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045538953,0.00022560709,0.007444257,0.001286332,0.0003467614,0.00035844158,0.003818159,0.083028324,0.031409185,0.27442122,0.010854157,0.5863521],"study_design_scores_gemma":[0.00009206739,0.0001816484,0.00195737,0.00019970047,0.00018762451,0.00040249032,0.00042323265,0.70754725,0.008330474,0.26418495,0.016364355,0.0001287331],"about_ca_topic_score_codex":0.0034911165,"about_ca_topic_score_gemma":0.0036796713,"teacher_disagreement_score":0.008384698,"about_ca_system_score_codex":0.0011931026,"about_ca_system_score_gemma":0.0014957523,"threshold_uncertainty_score":0.044343114},"labels":[],"label_agreement":null},{"id":"W3092435621","doi":"10.18653/v1/2020.emnlp-main.403","title":"On Losses for Modern Language Models","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Michael's Hospital; Vector Institute; University of Toronto","funders":"Canadian Institute for Advanced Research","keywords":"Computer science; Sentence; Benchmark (surveying); Task (project management); Baseline (sea); Natural language processing; Set (abstract data type); Context (archaeology); Artificial intelligence; Training set; Machine learning; Language model; Speech recognition; Programming language","score_opus":0.07732662353992366,"score_gpt":0.3006394215511403,"score_spread":0.22331279801121667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092435621","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02650611,0.0040892544,0.94561195,0.0041414145,0.0005925613,0.00014257546,0.0009807198,0.0048821885,0.013053189],"genre_scores_gemma":[0.5868233,0.0029376033,0.3482168,0.0033830716,0.0011670634,0.0008386393,0.004928016,0.0021798103,0.049525656],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99736387,0.0012887565,0.00013643764,0.0004378023,0.0005409461,0.00023211568],"domain_scores_gemma":[0.9913421,0.006807243,0.0001983568,0.00093166647,0.00055084453,0.00016990645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077320514,0.0024902984,0.0018078027,0.0011532326,0.00095129135,0.0026776944,0.0029694117,0.002597935,0.013208209],"category_scores_gemma":[0.020892905,0.00082347874,0.00094368256,0.0011144913,0.0015636577,0.0058202026,0.0032602483,0.0055489484,0.004864404],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043157663,0.00018863502,0.0011423298,0.00032352092,0.00015878437,0.00019791079,0.00010711303,0.6622263,0.0014902463,0.09334208,0.029864956,0.21052653],"study_design_scores_gemma":[0.000021185217,0.000037663038,0.00013272406,0.000027960326,0.000009741106,0.000035668883,0.000012885722,0.94880223,0.00049561716,0.04823252,0.0021828755,0.000008968649],"about_ca_topic_score_codex":0.005430872,"about_ca_topic_score_gemma":0.007982687,"teacher_disagreement_score":0.013208209,"about_ca_system_score_codex":0.0022950931,"about_ca_system_score_gemma":0.0017177869,"threshold_uncertainty_score":0.044185877},"labels":[],"label_agreement":null},{"id":"W3092539409","doi":"10.1016/j.nic.2020.08.001","title":"Review of Natural Language Processing in Radiology","year":2020,"lang":"en","type":"review","venue":"Neuroimaging Clinics of North America","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; McGill University","funders":"","keywords":"Artificial intelligence; Workflow; Natural language processing; Preprocessor; Deep learning; Computer science; Task (project management); Applications of artificial intelligence; Field (mathematics); Medicine","score_opus":0.044312118089355985,"score_gpt":0.36784689766127776,"score_spread":0.3235347795719218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092539409","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000046872316,0.999086,0.0001977443,0.000300558,0.00010295642,0.0000048740685,0.000026012352,0.0000083392915,0.00022656322],"genre_scores_gemma":[0.0004930935,0.99826556,0.00045476985,0.00035121368,0.000273684,0.000009002828,0.000048573962,0.0000028237382,0.00010125054],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992895,0.00022418833,0.0001707216,0.00012611688,0.00015331959,0.000036187204],"domain_scores_gemma":[0.99458647,0.004256489,0.00038167095,0.00009448812,0.00056048733,0.00012043691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021996337,0.0010113034,0.002027912,0.005208619,0.00030432735,0.0015145743,0.0013132866,0.0015949657,0.0042071575],"category_scores_gemma":[0.0073130205,0.00047217184,0.0011894986,0.004791161,0.00073014497,0.0023252864,0.0009601607,0.0016816678,0.0013368216],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007540269,0.000056092365,0.00028909105,0.049756657,0.0002837178,0.00011245613,0.00007008004,0.00038531746,0.00038704288,0.0017311668,0.03270513,0.9141478],"study_design_scores_gemma":[0.000077502664,0.00017699151,0.003358411,0.0481859,0.0017533946,0.0014913259,0.00016402011,0.0005837332,0.0005477953,0.0061731883,0.93740606,0.000081593585],"about_ca_topic_score_codex":0.0039138678,"about_ca_topic_score_gemma":0.0078006065,"teacher_disagreement_score":0.005208619,"about_ca_system_score_codex":0.0009090664,"about_ca_system_score_gemma":0.004165091,"threshold_uncertainty_score":0.014074385},"labels":[],"label_agreement":null},{"id":"W3092734951","doi":"10.18653/v1/2020.findings-emnlp.57","title":"From Language to Language-ish: How Brain-Like is an LSTM’s Representation of Nonsensical Language Stimuli?","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Language model; Representation (politics); Language identification; Natural language","score_opus":0.05953267168455737,"score_gpt":0.3398102721129692,"score_spread":0.2802776004284118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3092734951","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8249385,0.00102632,0.16281074,0.0027698302,0.0002166544,0.000039621307,0.0004713509,0.0006055658,0.0071214065],"genre_scores_gemma":[0.99162745,0.00023748778,0.0068876473,0.00020079313,0.00003376275,0.000012586439,0.00016597069,0.00006427258,0.00076990924],"study_design_codex":"bench_or_experimental","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99979204,0.00007568187,0.000007562354,0.00007889478,0.000023532783,0.00002239705],"domain_scores_gemma":[0.99932003,0.00034615945,0.0001173556,0.0000901972,0.00007288805,0.000053365005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006525372,0.00033060156,0.00027900285,0.00034309292,0.0001316368,0.001385476,0.00039233264,0.0006342403,0.0019202003],"category_scores_gemma":[0.006214757,0.00022056575,0.0004094552,0.000525873,0.00090633985,0.0036490664,0.00042491907,0.0009931258,0.00041314808],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008721538,0.00020262225,0.04169913,0.00085622177,0.00069940183,0.0007361585,0.0050825025,0.051494747,0.4930822,0.04804823,0.0052801403,0.35194653],"study_design_scores_gemma":[0.00006386058,0.00051252404,0.11492709,0.00017220006,0.0002637629,0.0007578762,0.0023970075,0.5749035,0.091094345,0.20926607,0.0054612383,0.00018053471],"about_ca_topic_score_codex":0.0011763917,"about_ca_topic_score_gemma":0.00084748864,"teacher_disagreement_score":0.0019202003,"about_ca_system_score_codex":0.00031098502,"about_ca_system_score_gemma":0.00020808041,"threshold_uncertainty_score":0.006423712},"labels":[],"label_agreement":null},{"id":"W3093493237","doi":"10.1145/3340531.3412746","title":"The Utility of Context When Extracting Entities From Legal Documents","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cisco Systems (Canada)","funders":"","keywords":"Computer science; Sentence; Context (archaeology); Natural language processing; Process (computing); Named-entity recognition; Artificial intelligence; Sequence (biology); Information retrieval; Layer (electronics); Legal document; Information extraction; Task (project management)","score_opus":0.04434289484593809,"score_gpt":0.25904431962161456,"score_spread":0.21470142477567647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093493237","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38434556,0.017014189,0.49894682,0.0043433495,0.0009857152,0.0008412704,0.013695299,0.02374157,0.056086242],"genre_scores_gemma":[0.6383814,0.00359059,0.34126925,0.00048571857,0.0004134721,0.00018903559,0.010197589,0.0006404174,0.0048325383],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986626,0.00042005617,0.00015099715,0.0004672566,0.00021417838,0.00008496336],"domain_scores_gemma":[0.99322903,0.0043295035,0.0005273756,0.0007843757,0.00093425816,0.00019534452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015390606,0.00080461433,0.00048763675,0.004760795,0.0010299049,0.0021429954,0.00074416935,0.0009828021,0.0037060906],"category_scores_gemma":[0.013215575,0.00040110268,0.00056947675,0.0029280612,0.00041254243,0.0065380074,0.0014018747,0.0011829495,0.0037459456],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054283446,0.00018847952,0.027342105,0.001105367,0.00011364,0.001294553,0.00213641,0.0076033343,0.031386517,0.00802013,0.022838559,0.89742815],"study_design_scores_gemma":[0.00015353349,0.0006984881,0.06726465,0.0012859254,0.00077236467,0.0058725746,0.004582636,0.37795743,0.13303053,0.05631211,0.35164368,0.00042608782],"about_ca_topic_score_codex":0.0051998505,"about_ca_topic_score_gemma":0.015620683,"teacher_disagreement_score":0.0051998505,"about_ca_system_score_codex":0.00049269776,"about_ca_system_score_gemma":0.0010183293,"threshold_uncertainty_score":0.012398124},"labels":[],"label_agreement":null},{"id":"W3093553144","doi":"10.18653/v1/2021.naacl-main.139","title":"UmlsBERT: Clinical Domain Knowledge Augmentation of Contextual Embeddings Using the Unified Medical Language System Metathesaurus","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Unified Medical Language System; Computer science; Natural language processing; Domain (mathematical analysis); Artificial intelligence; Embedding; Inference; Process (computing); Domain knowledge; Word (group theory); Word embedding; Named-entity recognition; ENCODE; Linguistics; Programming language","score_opus":0.07782002185679944,"score_gpt":0.38878477712195086,"score_spread":0.31096475526515144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093553144","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04875692,0.011475069,0.7954692,0.0037123295,0.0021045867,0.0010375808,0.05141476,0.06752836,0.018501032],"genre_scores_gemma":[0.23437372,0.0037948063,0.6370239,0.0009798888,0.00056417997,0.0009625035,0.110778175,0.0026005832,0.008922198],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986883,0.00048640269,0.00014823822,0.0003407815,0.00026144588,0.00007491707],"domain_scores_gemma":[0.9983045,0.00063344557,0.00012723118,0.0005025194,0.00034274976,0.000089557594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016502246,0.0011774999,0.00074700534,0.004437545,0.0006667544,0.0018299635,0.0012261476,0.0011803174,0.008923909],"category_scores_gemma":[0.007946474,0.00051731465,0.0012368155,0.002229273,0.0004066729,0.0028924209,0.00445679,0.001135722,0.0072613372],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061690365,0.00030691424,0.0033666557,0.0011052891,0.00038186097,0.00045843466,0.00070967665,0.007912529,0.00890311,0.008535305,0.15330999,0.81439334],"study_design_scores_gemma":[0.00080173556,0.0007190938,0.011791943,0.001463513,0.0010326168,0.0021788173,0.0017409556,0.32781044,0.03457243,0.09758093,0.5200172,0.00029026996],"about_ca_topic_score_codex":0.0050865486,"about_ca_topic_score_gemma":0.010213339,"teacher_disagreement_score":0.008923909,"about_ca_system_score_codex":0.0006447171,"about_ca_system_score_gemma":0.0021964847,"threshold_uncertainty_score":0.029853463},"labels":[],"label_agreement":null},{"id":"W3093562467","doi":"10.48550/arxiv.2010.12634","title":"Did You Ask a Good Question? A Cross-Domain Question Intention Classification Benchmark for Text-to-SQL","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Benchmark (surveying); Computer science; SQL; Task (project management); Set (abstract data type); Baseline (sea); Domain (mathematical analysis); Test set; Ask price; Null (SQL); Query by Example; Artificial intelligence; Information retrieval; Natural language processing; Data mining; Database; Programming language; Web search query; Search engine; Mathematics","score_opus":0.08798120581144701,"score_gpt":0.24458771559254633,"score_spread":0.1566065097810993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093562467","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75979364,0.0044505782,0.054821532,0.0027329281,0.0011120823,0.0017166316,0.1184292,0.034130286,0.02281314],"genre_scores_gemma":[0.54589224,0.0006510656,0.07759377,0.0011811298,0.0003391301,0.0013778327,0.35671735,0.0012797284,0.014967797],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9969023,0.0011706683,0.00033154764,0.00085667934,0.00049954414,0.0002392087],"domain_scores_gemma":[0.9878978,0.007677812,0.0005568959,0.0015594622,0.0015319788,0.00077591126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003947416,0.0026866896,0.0009337594,0.003557411,0.0010870752,0.0016661753,0.0019138221,0.0029406182,0.0074623246],"category_scores_gemma":[0.01651094,0.0003224488,0.0015738575,0.0021698652,0.00067652855,0.003470199,0.0025942149,0.0029720892,0.008228483],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005101387,0.007881488,0.0723671,0.0045095035,0.0009783925,0.0017474717,0.0023094185,0.048361313,0.026708858,0.0041966024,0.37893984,0.44689867],"study_design_scores_gemma":[0.00093808275,0.0029053076,0.11488089,0.000458877,0.0003852319,0.0023716039,0.0048086615,0.698428,0.047379564,0.014574034,0.11259817,0.000271671],"about_ca_topic_score_codex":0.010479331,"about_ca_topic_score_gemma":0.0132847605,"teacher_disagreement_score":0.010479331,"about_ca_system_score_codex":0.0011767817,"about_ca_system_score_gemma":0.0011051402,"threshold_uncertainty_score":0.024963975},"labels":[],"label_agreement":null},{"id":"W3093907107","doi":"10.1145/3340531.3412164","title":"Neural Relation Extraction on Wikipedia Tables for Augmenting Knowledge Graphs","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Knowledge graph; Relationship extraction; Table (database); Information retrieval; Task (project management); Information extraction; Question answering; Graph; Artificial neural network; Relation (database); Artificial intelligence; Knowledge extraction; Natural language processing; Machine learning; Data mining; Theoretical computer science","score_opus":0.061712156668974956,"score_gpt":0.29143019613185517,"score_spread":0.22971803946288022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093907107","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0997478,0.005374528,0.81839776,0.001339806,0.00067794864,0.00046907002,0.027837224,0.029365364,0.016790507],"genre_scores_gemma":[0.4154041,0.0025166436,0.50386745,0.0004854181,0.0002613308,0.00031373452,0.06813832,0.00078288093,0.008230156],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99935466,0.00011345886,0.0000492962,0.00029671236,0.00013817777,0.000047659963],"domain_scores_gemma":[0.99818194,0.0010508004,0.00012629842,0.00033970026,0.00024873888,0.00005260068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005471731,0.0011258525,0.0005607602,0.0055802003,0.0006723493,0.0012138804,0.0011787934,0.0009555248,0.004063442],"category_scores_gemma":[0.004329084,0.00043689524,0.0013133365,0.0042546396,0.0003583519,0.004077234,0.0010871068,0.0013324649,0.0026847585],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030201545,0.0004164412,0.009291821,0.0011279016,0.00029030532,0.001172249,0.00065958535,0.051208593,0.023376696,0.01996942,0.08767881,0.8045061],"study_design_scores_gemma":[0.000067903515,0.00015293364,0.007775086,0.0003406053,0.00037430212,0.0008272192,0.00054994377,0.7549972,0.03245613,0.089091524,0.11327752,0.00008957299],"about_ca_topic_score_codex":0.009091173,"about_ca_topic_score_gemma":0.026758049,"teacher_disagreement_score":0.009091173,"about_ca_system_score_codex":0.0007200773,"about_ca_system_score_gemma":0.0010895,"threshold_uncertainty_score":0.01807648},"labels":[],"label_agreement":null},{"id":"W3093977588","doi":"10.48550/arxiv.2010.11137","title":"Multi-Domain Dialogue State Tracking based on State Graph","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"State (computer science); Computer science; Graph; Domain (mathematical analysis); Tracking (education); Artificial intelligence; Theoretical computer science; Algorithm; Mathematics; Psychology","score_opus":0.11301556467133135,"score_gpt":0.19827144409052608,"score_spread":0.08525587941919473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3093977588","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033678945,0.00065004925,0.9571346,0.00017264752,0.00008250802,0.00008715422,0.00083343574,0.0056671496,0.001693515],"genre_scores_gemma":[0.69025,0.00041582237,0.30096716,0.00017106705,0.000068962865,0.00013497201,0.003583844,0.00047368818,0.003934471],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99868065,0.00033850697,0.000053082276,0.00064501405,0.00019741521,0.00008534854],"domain_scores_gemma":[0.9983072,0.0010182329,0.0001392176,0.00026305977,0.00019557668,0.00007673815],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009912031,0.000987036,0.00079102826,0.0016241102,0.00050153246,0.0010036492,0.0011262777,0.0008623608,0.0017016117],"category_scores_gemma":[0.0034464642,0.00038310417,0.0009491253,0.0012971035,0.0005625788,0.0030615258,0.0014898068,0.0013795571,0.0010383843],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007623713,0.00026361636,0.005006643,0.00049925153,0.00020274431,0.00049393723,0.0015583514,0.123422295,0.037437618,0.021687524,0.008126901,0.8005387],"study_design_scores_gemma":[0.000015497435,0.00006458337,0.001072877,0.000021445057,0.000052867726,0.00011951493,0.00012531245,0.96711785,0.010511532,0.016847549,0.0040209563,0.000030066183],"about_ca_topic_score_codex":0.007507972,"about_ca_topic_score_gemma":0.0075312965,"teacher_disagreement_score":0.007507972,"about_ca_system_score_codex":0.0007396977,"about_ca_system_score_gemma":0.0007757243,"threshold_uncertainty_score":0.014928579},"labels":[],"label_agreement":null},{"id":"W3094834348","doi":"10.2196/23375","title":"The 2019 n2c2/OHNLP Track on Clinical Semantic Textual Similarity: Overview","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; U.S. National Library of Medicine","keywords":"Computer science; Automatic summarization; Natural language processing; Information retrieval; Semantic similarity; Artificial intelligence; Task (project management); Unified Medical Language System; Semantics (computer science); Sentence; Similarity (geometry); Set (abstract data type); Semantic computing; Semantic Web","score_opus":0.09220643652773049,"score_gpt":0.3689164150421732,"score_spread":0.27670997851444273,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094834348","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02911266,0.011075085,0.12287691,0.066518456,0.015044444,0.0074630533,0.6505544,0.05209559,0.0452593],"genre_scores_gemma":[0.009165428,0.0015281081,0.062432937,0.005842403,0.001231439,0.0035494578,0.8999338,0.0033822618,0.012934141],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96973675,0.008103461,0.0025952822,0.0051384624,0.0122161275,0.0022100084],"domain_scores_gemma":[0.8889642,0.029684035,0.0038215513,0.017994527,0.043927494,0.015608144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.042028327,0.0035960362,0.0035172396,0.013243121,0.005255091,0.009386616,0.009724856,0.0060153822,0.032504614],"category_scores_gemma":[0.074837364,0.0018089492,0.0035181118,0.010972812,0.0024945957,0.016420467,0.018605025,0.009110782,0.03593402],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033312957,0.00033807577,0.0024321743,0.0009199041,0.000107734,0.0001501985,0.00028377818,0.0010829783,0.0026909765,0.0015832966,0.9386088,0.051468916],"study_design_scores_gemma":[0.000497443,0.0004507262,0.010152561,0.0005985545,0.00012132571,0.00038293021,0.00078681804,0.01385841,0.005108189,0.007181377,0.96067494,0.00018670931],"about_ca_topic_score_codex":0.03949036,"about_ca_topic_score_gemma":0.06630088,"teacher_disagreement_score":0.042028327,"about_ca_system_score_codex":0.005933921,"about_ca_system_score_gemma":0.019684248,"threshold_uncertainty_score":0.22226965},"labels":[],"label_agreement":null},{"id":"W3095642204","doi":"10.2196/19735","title":"Measurement of Semantic Textual Similarity in Clinical Texts: Comparison of Transformer-Based Models","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Chronic Disease Prevention and Health Promotion; Clinical and Translational Science Institute, University of Florida; National Center for Advancing Translational Sciences; University of Florida; National Cancer Institute; National Institutes of Health; Nvidia; Centers for Disease Control and Prevention; National Institute on Aging; Patient-Centered Outcomes Research Institute","keywords":"Computer science; Natural language processing; Semantic similarity; Artificial intelligence; Similarity (geometry); Transformer; Domain (mathematical analysis); Information retrieval; Biomedical text mining; Text mining","score_opus":0.1322531638745227,"score_gpt":0.36620898419324677,"score_spread":0.23395582031872406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095642204","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6021012,0.008245749,0.3559785,0.0032348563,0.00079218735,0.0006851774,0.0052817953,0.0076275477,0.016053028],"genre_scores_gemma":[0.9419789,0.0013561499,0.0470772,0.00032505472,0.00015547867,0.00017235248,0.005590871,0.00020907342,0.0031349487],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99785054,0.0009098941,0.00018540511,0.00051006227,0.00041305626,0.00013104617],"domain_scores_gemma":[0.9908325,0.006475816,0.00043305685,0.00067722716,0.0012641681,0.00031728685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050161756,0.0013128583,0.0007659723,0.00393894,0.00042048964,0.0016212728,0.001384937,0.001176812,0.0021593887],"category_scores_gemma":[0.017373173,0.00032930964,0.0012796529,0.0018956356,0.0006506903,0.0033665935,0.0015557397,0.0017659282,0.0014691547],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002330955,0.0009933824,0.043418285,0.00068809197,0.00079174014,0.0003579667,0.0008116547,0.30747712,0.005552152,0.009149016,0.015602371,0.6128273],"study_design_scores_gemma":[0.000029277957,0.00026833842,0.003564548,0.000047560763,0.00008953998,0.00013295427,0.00016415547,0.98732126,0.0018899036,0.005370048,0.0010936144,0.000028789293],"about_ca_topic_score_codex":0.013078058,"about_ca_topic_score_gemma":0.0153818065,"teacher_disagreement_score":0.013078058,"about_ca_system_score_codex":0.0023340276,"about_ca_system_score_gemma":0.0016683993,"threshold_uncertainty_score":0.026528418},"labels":[],"label_agreement":null},{"id":"W3096266342","doi":"10.1007/978-3-030-61377-8_28","title":"BERTimbau: Pretrained BERT Models for Brazilian Portuguese","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":598,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Transformer; Language model; Sentence; Encoder; Portuguese; Transfer of learning; Textual entailment; Named-entity recognition; Logical consequence; Linguistics","score_opus":0.029260968444769126,"score_gpt":0.25291092292700496,"score_spread":0.22364995448223582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096266342","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08035384,0.0032864816,0.7756815,0.0046134396,0.0008435063,0.0001623093,0.010073113,0.004093669,0.12089214],"genre_scores_gemma":[0.74947745,0.0019855404,0.12897764,0.0004105178,0.00028640288,0.0002636437,0.0074339984,0.0019270743,0.10923769],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99972814,0.00009838288,0.000014078021,0.000058155,0.00004917988,0.00005199624],"domain_scores_gemma":[0.9993741,0.0003527912,0.000039035698,0.00008598504,0.00010226373,0.000045889217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009027429,0.0012110532,0.0008981635,0.00052037416,0.0008505856,0.0020344045,0.0017743718,0.0015011287,0.015197346],"category_scores_gemma":[0.004401272,0.0006696854,0.0012501123,0.00081597996,0.0004539629,0.0022015632,0.0010675157,0.001811847,0.0024446459],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017798285,0.000063366126,0.0010495366,0.00013589457,0.000052358308,0.0001755508,0.00022432164,0.7455327,0.00060861214,0.17086875,0.024075154,0.05703569],"study_design_scores_gemma":[0.0000119183605,0.0000078567955,0.00021062231,0.000025181724,0.00001329347,0.000022709859,0.000030920128,0.9444594,0.00021063977,0.04464785,0.010348361,0.000011318339],"about_ca_topic_score_codex":0.07336289,"about_ca_topic_score_gemma":0.08831627,"teacher_disagreement_score":0.07336289,"about_ca_system_score_codex":0.0019788663,"about_ca_system_score_gemma":0.001749722,"threshold_uncertainty_score":0.14587176},"labels":[],"label_agreement":null},{"id":"W3096316167","doi":"10.48550/arxiv.2011.02944","title":"Learning Efficient Task-Specific Meta-Embeddings with Word Prisms","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Word (group theory); Computer science; Embedding; Inference; Set (abstract data type); Task (project management); Natural language processing; Artificial intelligence; Word embedding; Context (archaeology); Simple (philosophy); Meta learning (computer science); Space (punctuation); Mathematics","score_opus":0.10821251177931261,"score_gpt":0.18441847499740743,"score_spread":0.07620596321809482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096316167","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038473096,0.00096929393,0.94584006,0.00038129126,0.0001814028,0.00011874368,0.0012759125,0.011157288,0.0016029526],"genre_scores_gemma":[0.45246872,0.00082316383,0.5301941,0.00032999314,0.00016629216,0.00045586543,0.009049787,0.0012807797,0.005231258],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986085,0.00047255453,0.0001418483,0.00044646833,0.00022222957,0.000108350265],"domain_scores_gemma":[0.99719524,0.0009877845,0.00022727753,0.0010706441,0.00039795318,0.00012107142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021098936,0.0023445943,0.0011674678,0.0013564315,0.00034541864,0.0016448379,0.002383025,0.001582825,0.0027738276],"category_scores_gemma":[0.0076710964,0.00090848317,0.0018608448,0.0020894671,0.000781886,0.007922399,0.0037391966,0.0031730325,0.0024914],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006280782,0.00046800502,0.0049894406,0.00079281133,0.00042847247,0.0002648789,0.00044958264,0.21777296,0.019287212,0.01938953,0.020273669,0.7152554],"study_design_scores_gemma":[0.00007485211,0.00016461426,0.0005048619,0.00004080976,0.000069346825,0.00012633565,0.000104303435,0.9526964,0.008684302,0.03354448,0.003956234,0.000033408607],"about_ca_topic_score_codex":0.0015378433,"about_ca_topic_score_gemma":0.0033255543,"teacher_disagreement_score":0.0027738276,"about_ca_system_score_codex":0.00061460724,"about_ca_system_score_gemma":0.0014187895,"threshold_uncertainty_score":0.011158347},"labels":[],"label_agreement":null},{"id":"W3096437928","doi":"","title":"Leveraging a Domain Ontology in (Neural) Learning from Heterogeneous Data.","year":2020,"lang":"en","type":"article","venue":"Conference on Information and Knowledge Management","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Ontology; Domain (mathematical analysis); Artificial intelligence; Machine learning; Data science","score_opus":0.0830432757915092,"score_gpt":0.27142507716249303,"score_spread":0.18838180137098381,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096437928","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05151501,0.0023929416,0.93783116,0.0014674845,0.0002297604,0.0001977248,0.0013078643,0.001155032,0.0039030965],"genre_scores_gemma":[0.6210547,0.001758212,0.36810863,0.00063457375,0.0002272694,0.0002934856,0.0050778682,0.0001873281,0.0026579166],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978536,0.0010051614,0.00020379477,0.0004255906,0.00037817055,0.00013371359],"domain_scores_gemma":[0.995445,0.002915408,0.000249848,0.0006680631,0.0005715631,0.00015008882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038805797,0.0004909404,0.0006040323,0.0027112826,0.00081423594,0.0021670156,0.0014037911,0.0012222083,0.0010771885],"category_scores_gemma":[0.011874502,0.00031988721,0.0011735947,0.0032174857,0.00065232086,0.0058659893,0.0025770257,0.0018955157,0.0005664442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045094363,0.0008286532,0.014037813,0.0007001267,0.0006716189,0.0005959266,0.000995139,0.11904401,0.009522957,0.03971623,0.017128948,0.7963076],"study_design_scores_gemma":[0.000023532635,0.00006013626,0.0017662266,0.00011411139,0.00017986,0.00014502993,0.00037184003,0.9128834,0.003496542,0.07230556,0.008624908,0.000028737146],"about_ca_topic_score_codex":0.008291276,"about_ca_topic_score_gemma":0.014378333,"teacher_disagreement_score":0.008291276,"about_ca_system_score_codex":0.0009544184,"about_ca_system_score_gemma":0.0015878562,"threshold_uncertainty_score":0.020522714},"labels":[],"label_agreement":null},{"id":"W3096542808","doi":"10.1145/3402884","title":"Condition-Transforming Variational Autoencoder for Generating Diverse Short Text Conversations","year":2020,"lang":"en","type":"article","venue":"ACM Transactions on Asian and Low-Resource Language Information Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China","keywords":"Latent variable; Autoencoder; Sequence (biology); Dependency (UML); Conditional probability distribution; Computer science; Artificial intelligence; Gaussian; Latent variable model; Pattern recognition (psychology); Transformation (genetics); Variable (mathematics); Multivariate normal distribution; Multivariate statistics; Algorithm; Mathematics; Statistics; Machine learning","score_opus":0.01530751449655362,"score_gpt":0.2482199830368645,"score_spread":0.23291246854031086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096542808","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01301014,0.00025018444,0.98500174,0.00012241577,0.00005036099,0.000048279417,0.00016132185,0.0005361019,0.00081939634],"genre_scores_gemma":[0.52839667,0.000548618,0.4601647,0.00047108947,0.00012831828,0.0004368111,0.0019444752,0.00040828227,0.0075010182],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99933845,0.00026116986,0.000034205408,0.00020158292,0.00010609572,0.000058441172],"domain_scores_gemma":[0.99876,0.0008895738,0.00005521728,0.00008757309,0.00016073903,0.000046852998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012016314,0.0010104914,0.0007156623,0.00036207022,0.00030556647,0.00047876526,0.0010076895,0.0008556173,0.0026686518],"category_scores_gemma":[0.004075896,0.00047554707,0.0009804509,0.00041839547,0.0005427188,0.0012396189,0.000976488,0.0019260176,0.0009880954],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003651009,0.00013563519,0.00158494,0.00020696467,0.00014379506,0.00015559077,0.00037150382,0.6071214,0.02582523,0.020832255,0.0044273203,0.3388302],"study_design_scores_gemma":[0.0000063375123,0.000019101217,0.00010845392,0.0000051689817,0.0000070058277,0.000014639142,0.000010417191,0.99399996,0.0020897086,0.0032658956,0.00046638955,0.00000705167],"about_ca_topic_score_codex":0.004023741,"about_ca_topic_score_gemma":0.0065778526,"teacher_disagreement_score":0.004023741,"about_ca_system_score_codex":0.0005058578,"about_ca_system_score_gemma":0.00088338624,"threshold_uncertainty_score":0.008927524},"labels":[],"label_agreement":null},{"id":"W3096783161","doi":"10.18653/v1/2020.coling-main.495","title":"WSL-DS: Weakly Supervised Learning with Distant Supervision for Query Focused Multi-Document Abstractive Summarization","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Automatic summarization; Computer science; Leverage (statistics); Task (project management); Multi-document summarization; Set (abstract data type); Information retrieval; Sentence; Artificial intelligence; Training set; Similarity (geometry); Natural language processing","score_opus":0.04232295421965686,"score_gpt":0.27204804479948835,"score_spread":0.22972509057983148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096783161","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022614822,0.0016722936,0.9660692,0.0004194752,0.00009588546,0.00020545516,0.0009088635,0.006963964,0.0010500676],"genre_scores_gemma":[0.40894237,0.0009050235,0.56754833,0.0009083886,0.00045876042,0.0006476507,0.01221306,0.00068893057,0.0076874853],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982893,0.0006722729,0.0001241261,0.0005087702,0.0003197814,0.000085783395],"domain_scores_gemma":[0.99749136,0.0011212924,0.0002422495,0.00046599304,0.0005519589,0.00012717667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002311293,0.001487181,0.0015998245,0.001709827,0.00061777915,0.0010657255,0.0025488846,0.0014715483,0.0018382767],"category_scores_gemma":[0.0056035733,0.00043975384,0.00096518465,0.0016630215,0.0006358205,0.0029250362,0.0015548377,0.0019409187,0.0013197372],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000576513,0.00070951384,0.0023708711,0.0006726215,0.0003245618,0.0002398499,0.00048455264,0.14998916,0.019663496,0.005952295,0.030572701,0.7884439],"study_design_scores_gemma":[0.000057803314,0.00022464455,0.00045723372,0.000016064632,0.000038726266,0.000057111632,0.00006993481,0.981834,0.004987947,0.008870274,0.0033663877,0.00001990999],"about_ca_topic_score_codex":0.0036551512,"about_ca_topic_score_gemma":0.0065698833,"teacher_disagreement_score":0.0036551512,"about_ca_system_score_codex":0.00091254985,"about_ca_system_score_gemma":0.0013481446,"threshold_uncertainty_score":0.0122234225},"labels":[],"label_agreement":null},{"id":"W3096831026","doi":"10.18653/v1/2020.coling-main.208","title":"Automatic Detection of Machine Generated Text: A Critical Survey","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Computer science; Generative grammar; Data science; Natural language processing; Artificial intelligence; Product (mathematics); Machine learning; Information retrieval","score_opus":0.058270004379648195,"score_gpt":0.3015808177752705,"score_spread":0.24331081339562233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096831026","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016000586,0.55318123,0.3894619,0.012186527,0.0025605082,0.000384842,0.0031747576,0.0074235764,0.015626019],"genre_scores_gemma":[0.18857004,0.39489904,0.36617312,0.007435936,0.009149012,0.00069334,0.014039503,0.0028052807,0.016234733],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9896818,0.0023824363,0.00073311006,0.0025656854,0.0043634395,0.0002734726],"domain_scores_gemma":[0.9334529,0.051257912,0.0019045657,0.0043662665,0.008504884,0.00051344873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007714324,0.0021103546,0.0030830957,0.01240316,0.00109491,0.0053363545,0.0055008917,0.0038053202,0.004569533],"category_scores_gemma":[0.034886807,0.0017774157,0.0015061729,0.005097547,0.002359195,0.012660562,0.0022190814,0.003125046,0.007721825],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009579744,0.0001239984,0.0049517197,0.0038074655,0.00013229404,0.00015803488,0.00039072716,0.005176946,0.0030131787,0.015673008,0.05828948,0.9081874],"study_design_scores_gemma":[0.00005481574,0.0004114381,0.011913201,0.0051278556,0.0003549447,0.0042320285,0.0015557928,0.21253178,0.040017508,0.079104565,0.64424545,0.00045061242],"about_ca_topic_score_codex":0.0022874312,"about_ca_topic_score_gemma":0.0018563529,"teacher_disagreement_score":0.01240316,"about_ca_system_score_codex":0.0019789885,"about_ca_system_score_gemma":0.0019084512,"threshold_uncertainty_score":0.04079777},"labels":[],"label_agreement":null},{"id":"W3096959275","doi":"10.1007/978-3-030-63128-4_58","title":"Learning Reddit User Reputation Using Graphical Attention Networks","year":2020,"lang":"en","type":"book-chapter","venue":"Advances in intelligent systems and computing","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Reputation; Embedding; Predictive power; Focus (optics); Task (project management); Graph; Machine learning; Set (abstract data type); Artificial intelligence; Attention network; Feature (linguistics); Feature engineering; Simple (philosophy); Deep learning; Feature learning; Graphical model; Human–computer interaction; Theoretical computer science","score_opus":0.025281091859025755,"score_gpt":0.26914745542601975,"score_spread":0.243866363566994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096959275","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11890345,0.00363463,0.8622921,0.0015932964,0.00050160976,0.0001104388,0.0005855075,0.005907384,0.006471549],"genre_scores_gemma":[0.9242254,0.0005630465,0.06346449,0.00035152189,0.00047384293,0.00006471421,0.0008479109,0.0003067783,0.009702276],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982515,0.00059213286,0.000080839534,0.0006024102,0.00026092512,0.00021220042],"domain_scores_gemma":[0.98981744,0.0074592694,0.00061372964,0.00093052455,0.0008213166,0.00035762804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037002554,0.0014497484,0.0022968221,0.0026506807,0.00070189545,0.0017163312,0.0025311355,0.0022653728,0.0036609734],"category_scores_gemma":[0.013826591,0.00083145534,0.0010532993,0.0018713884,0.0008736074,0.0038610739,0.0017755133,0.0029635655,0.0018673126],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012232336,0.0007273384,0.016071722,0.00025045552,0.0005282583,0.00023395612,0.00025432828,0.326358,0.0053519215,0.015894959,0.02710213,0.60600364],"study_design_scores_gemma":[0.000009500457,0.000038230726,0.0005642736,0.000006187627,0.000028975352,0.000028141005,0.0000073680653,0.9939021,0.00041654217,0.004719213,0.0002699323,0.000009534749],"about_ca_topic_score_codex":0.007314288,"about_ca_topic_score_gemma":0.011985371,"teacher_disagreement_score":0.007314288,"about_ca_system_score_codex":0.0017274569,"about_ca_system_score_gemma":0.00068863336,"threshold_uncertainty_score":0.01956904},"labels":[],"label_agreement":null},{"id":"W3097602436","doi":"10.18653/v1/2020.emnlp-main.648","title":"Multi-XScience: A Large-scale Dataset for Extreme Multi-document Summarization of Scientific Articles","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal; McGill University","funders":"Institut de Valorisation des Données; Compute Canada; Canadian Institute for Advanced Research","keywords":"Automatic summarization; Computer science; Task (project management); Scale (ratio); Multi-document summarization; Information retrieval; Data science; Natural language processing; Artificial intelligence; Engineering","score_opus":0.12750749254199562,"score_gpt":0.3234923945206476,"score_spread":0.19598490197865198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3097602436","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08322953,0.0068542706,0.10414397,0.003700434,0.0018240412,0.0017947847,0.7504591,0.034663133,0.013330732],"genre_scores_gemma":[0.04990854,0.0009805695,0.13231906,0.00044748097,0.0004279871,0.001113814,0.8105431,0.000622207,0.0036373108],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978161,0.0005311197,0.00040712574,0.0005089805,0.0006287255,0.00010785063],"domain_scores_gemma":[0.9936045,0.0022770534,0.00077141664,0.0013483609,0.0014848005,0.00051385834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020286755,0.0011215066,0.00075822405,0.006439524,0.0009840131,0.001635022,0.0015429166,0.0016577581,0.0045758924],"category_scores_gemma":[0.009728466,0.00029480137,0.0011724443,0.0056420425,0.00053102203,0.0024351268,0.0019081142,0.0015965072,0.0051273275],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012035496,0.0008321093,0.013794919,0.0064680013,0.0006150287,0.0007082864,0.0011270926,0.012502994,0.054450985,0.0076885927,0.5865784,0.31403005],"study_design_scores_gemma":[0.00067180576,0.00095141394,0.0511893,0.0004889155,0.0003581628,0.0009666763,0.0012091639,0.07089893,0.051232614,0.016720856,0.8049521,0.00035997943],"about_ca_topic_score_codex":0.00263149,"about_ca_topic_score_gemma":0.008596601,"teacher_disagreement_score":0.006439524,"about_ca_system_score_codex":0.0007917975,"about_ca_system_score_gemma":0.0019850377,"threshold_uncertainty_score":0.015307844},"labels":[],"label_agreement":null},{"id":"W3097657226","doi":"10.23977/jaip.2020.030107","title":"The Progess That Natural Language Processing Has Made Towards Human-level AI","year":2020,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Practice","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Human language; Context (archaeology); Natural language understanding; Natural (archaeology); Language technology; Natural language; Linguistics; History; Comprehension approach; Philosophy","score_opus":0.20598678695038364,"score_gpt":0.39334098531349554,"score_spread":0.1873541983631119,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3097657226","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018777767,0.13394791,0.3548681,0.225081,0.005006437,0.00014282606,0.00048751288,0.0013448081,0.2603437],"genre_scores_gemma":[0.36834553,0.15543458,0.3810594,0.03417202,0.01381455,0.00026751735,0.00069972966,0.00085722946,0.04534936],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9941479,0.0029526511,0.0002209982,0.00076082186,0.0017197763,0.00019793457],"domain_scores_gemma":[0.9705244,0.019440677,0.001204744,0.0047062174,0.0033113672,0.0008126832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011626828,0.00074869726,0.0005775201,0.0023855474,0.002112043,0.010822037,0.0013010622,0.0029367707,0.010896986],"category_scores_gemma":[0.019923914,0.0005110043,0.00068574195,0.002669128,0.014995903,0.019217702,0.004596985,0.006248216,0.004409213],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007660234,0.00007690522,0.002295748,0.0014554706,0.000062116036,0.000103453545,0.0021222704,0.002283836,0.0018602266,0.8497337,0.015136573,0.12479313],"study_design_scores_gemma":[0.0000106021125,0.00010959486,0.0015640236,0.00075097475,0.000027180546,0.0003607313,0.0010159006,0.0043462734,0.0016662423,0.49870878,0.49137115,0.00006849618],"about_ca_topic_score_codex":0.0034279402,"about_ca_topic_score_gemma":0.0029291513,"teacher_disagreement_score":0.011626828,"about_ca_system_score_codex":0.0022712962,"about_ca_system_score_gemma":0.002092015,"threshold_uncertainty_score":0.061489284},"labels":[],"label_agreement":null},{"id":"W3098123193","doi":"10.18653/v1/2020.codi-1.13","title":"Do We Really Need That Many Parameters In Transformer For Extractive Summarization? Discourse Can Help !","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Huawei Technologies","keywords":"Automatic summarization; Transformer; Computer science; Artificial intelligence; Natural language processing; Sentence; Language model; Machine learning","score_opus":0.04987185070598387,"score_gpt":0.2862067534909244,"score_spread":0.23633490278494051,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098123193","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016890794,0.0043771146,0.92513204,0.021343842,0.0010104554,0.00018990676,0.0008115824,0.016732767,0.01351143],"genre_scores_gemma":[0.35952044,0.0032143868,0.6069119,0.005511379,0.0011106175,0.00030809033,0.0018920902,0.006916904,0.014614214],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975781,0.0012018688,0.00015382936,0.00055192324,0.00032952594,0.0001846794],"domain_scores_gemma":[0.98976,0.0056642056,0.00044069102,0.0024856469,0.001210183,0.00043927637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005102359,0.0019609951,0.0015790075,0.0013460917,0.0014290694,0.0038671584,0.0023062965,0.0026839261,0.02281682],"category_scores_gemma":[0.029523546,0.0010440666,0.0011247396,0.0012503744,0.00196873,0.02774908,0.0033167405,0.005272944,0.014617579],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012864423,0.00022715569,0.0028843768,0.0013140282,0.00023233457,0.0003179995,0.0035043077,0.009126462,0.0383521,0.099186145,0.05099892,0.79256976],"study_design_scores_gemma":[0.00030901798,0.0005044207,0.0022995935,0.0008842827,0.00058998045,0.0008577961,0.0037827624,0.1310708,0.05400408,0.55268365,0.25264546,0.0003682244],"about_ca_topic_score_codex":0.0028729595,"about_ca_topic_score_gemma":0.006120008,"teacher_disagreement_score":0.02281682,"about_ca_system_score_codex":0.0011808532,"about_ca_system_score_gemma":0.001492767,"threshold_uncertainty_score":0.07632983},"labels":[],"label_agreement":null},{"id":"W3098135292","doi":"10.18653/v1/2020.emnlp-main.427","title":"Help! Need Advice on Identifying Advice","year":2020,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"University of Texas at Austin; National Science Foundation","keywords":"Advice (programming); Pragmatics; Computer science; Variety (cybernetics); GRASP; Psychology; Artificial intelligence; Linguistics","score_opus":0.08848159657507164,"score_gpt":0.285614763310375,"score_spread":0.19713316673530334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098135292","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5723595,0.008513544,0.14752196,0.018650675,0.0017583768,0.0012158479,0.10377786,0.040214635,0.105987586],"genre_scores_gemma":[0.7126443,0.0022256891,0.12612692,0.0029513063,0.0005939503,0.00046886186,0.11448586,0.0021683085,0.038334895],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99754876,0.0009874065,0.00014605669,0.0006015173,0.0005543432,0.0001618599],"domain_scores_gemma":[0.9844848,0.010178247,0.0007447211,0.0015143055,0.0024690365,0.0006088679],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002171968,0.0009159209,0.0004492831,0.0019835664,0.0011838693,0.0018882703,0.0008739271,0.0017997717,0.00944905],"category_scores_gemma":[0.024507772,0.00045237536,0.000649866,0.0012972888,0.0006529814,0.0041336385,0.0015257543,0.0019219045,0.009496638],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014029897,0.0005610014,0.12347247,0.0020608183,0.00011797059,0.00090982724,0.006619223,0.009982688,0.009944806,0.01083562,0.38950345,0.4445891],"study_design_scores_gemma":[0.00020966562,0.00039791683,0.10901714,0.0007161806,0.00018056769,0.0017876705,0.003893073,0.2801016,0.015845953,0.025703622,0.5619542,0.00019239854],"about_ca_topic_score_codex":0.025695695,"about_ca_topic_score_gemma":0.03874489,"teacher_disagreement_score":0.025695695,"about_ca_system_score_codex":0.0008635058,"about_ca_system_score_gemma":0.0013952744,"threshold_uncertainty_score":0.051092207},"labels":[],"label_agreement":null},{"id":"W3098136301","doi":"10.18653/v1/2020.emnlp-main.748","title":"On Extractive and Abstractive Neural Document Summarization with Transformer Language Models","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":187,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; Université de Montréal; Polytechnique Montréal","funders":"","keywords":"Automatic summarization; Computer science; Transformer; Language model; Natural language processing; Artificial intelligence; Fluency; Question answering; Machine learning; Artificial neural network; Information retrieval; Linguistics","score_opus":0.01696996314490283,"score_gpt":0.2343981941818167,"score_spread":0.2174282310369139,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098136301","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024481427,0.0018494395,0.96264017,0.0004921124,0.000116918825,0.00015152051,0.0006733978,0.006778708,0.0028163174],"genre_scores_gemma":[0.34601468,0.0020046849,0.6344753,0.00045556694,0.00037581642,0.00037244966,0.0047546946,0.00071135856,0.010835462],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99919325,0.00025805668,0.00008476919,0.00017799195,0.00023306829,0.000052827854],"domain_scores_gemma":[0.9969439,0.001638773,0.00029531482,0.00041142348,0.0006089617,0.00010163352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014779592,0.0013386862,0.00074775965,0.0018513968,0.0004062157,0.0014073893,0.0011686326,0.0008829529,0.002488843],"category_scores_gemma":[0.006973643,0.0003691692,0.00086358027,0.0015969564,0.00053902785,0.0032181311,0.0011289819,0.0013616749,0.0018939062],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042638337,0.00021047263,0.0011213529,0.00055153377,0.00018668691,0.00021371941,0.00050097064,0.12165071,0.03937862,0.012643289,0.008895063,0.81422114],"study_design_scores_gemma":[0.00006369646,0.00040018524,0.00068939093,0.00005228267,0.00012809591,0.00015598827,0.00016213176,0.9361052,0.029415205,0.022585383,0.010187527,0.000054899097],"about_ca_topic_score_codex":0.0040910314,"about_ca_topic_score_gemma":0.0084882965,"teacher_disagreement_score":0.0040910314,"about_ca_system_score_codex":0.0008537527,"about_ca_system_score_gemma":0.0011570002,"threshold_uncertainty_score":0.008326054},"labels":[],"label_agreement":null},{"id":"W3098382480","doi":"10.1007/978-3-030-63591-6_63","title":"Utilizing Bidirectional Encoder Representations from Transformers for Answer Selection","year":2021,"lang":"en","type":"preprint","venue":"Springer proceedings in mathematics & statistics","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Transformer; Language model; Question answering; Encoder; Selection (genetic algorithm); Sentence; Artificial intelligence; Natural language processing; Task (project management); Machine learning; Voltage","score_opus":0.04253244669200808,"score_gpt":0.306499181081527,"score_spread":0.26396673438951895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098382480","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020087773,0.0004285005,0.9679501,0.0005735507,0.0001966038,0.0001316727,0.0016192435,0.0058499523,0.0031626348],"genre_scores_gemma":[0.6066667,0.0007239473,0.37745348,0.00038206388,0.00036759375,0.00035573935,0.006694316,0.0011534921,0.0062027257],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99868053,0.00047181046,0.0000954887,0.0002753377,0.00030154112,0.00017522673],"domain_scores_gemma":[0.9952349,0.0030455503,0.0001558799,0.00069010083,0.0007072548,0.00016634207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017509642,0.00090031844,0.0010218102,0.0018551226,0.0005288039,0.0020826962,0.0013627777,0.0011723081,0.010400655],"category_scores_gemma":[0.011093294,0.0005043648,0.0009363997,0.0018995246,0.00062064116,0.0048739766,0.0023031752,0.0018775656,0.003812177],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012658013,0.0003230164,0.0018779348,0.00044556352,0.0001085605,0.0002483963,0.00046151443,0.025761947,0.017805425,0.14589715,0.031980135,0.77382463],"study_design_scores_gemma":[0.000109308756,0.00013205239,0.0003333063,0.00007210454,0.00010763502,0.00011197994,0.00017392624,0.69680893,0.014906879,0.27834216,0.008863607,0.00003804044],"about_ca_topic_score_codex":0.002497115,"about_ca_topic_score_gemma":0.004764995,"teacher_disagreement_score":0.010400655,"about_ca_system_score_codex":0.00074982,"about_ca_system_score_gemma":0.001714524,"threshold_uncertainty_score":0.034793675},"labels":[],"label_agreement":null},{"id":"W3098396746","doi":"10.18653/v1/2020.emnlp-main.551","title":"Distilling Structured Knowledge for Text-Based Relational Reasoning","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); Canadian Institute for Advanced Research; McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Canadian Institute for Advanced Research; Compute Canada; Microsoft Research","keywords":"Computer science; Artificial intelligence; Natural language processing; Benchmark (surveying); Task (project management); Statistical relational learning; Machine learning; Relational database; Information retrieval","score_opus":0.04633304072299972,"score_gpt":0.2667868192989307,"score_spread":0.22045377857593101,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098396746","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033104766,0.0013428308,0.94186187,0.0014783568,0.0001574978,0.00025440333,0.0043498673,0.012580536,0.0048698387],"genre_scores_gemma":[0.35245836,0.0010718227,0.62272793,0.0006689418,0.00016863176,0.00026916724,0.018531822,0.0005448524,0.0035585372],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99863607,0.00042356137,0.00010853713,0.00047377692,0.00028052265,0.00007754883],"domain_scores_gemma":[0.99648523,0.002347902,0.00020121604,0.0005661077,0.00031664607,0.00008285342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018412557,0.0014730288,0.000705339,0.0027521518,0.0006120533,0.0017581359,0.002781794,0.0016281061,0.00611332],"category_scores_gemma":[0.009024005,0.00050552184,0.0020441962,0.0021579333,0.0008469356,0.0083086565,0.0025559417,0.0028096936,0.0032501782],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003744748,0.00046216973,0.0021967154,0.0012763337,0.00024353995,0.0006993085,0.00093130046,0.23732421,0.0214895,0.04545273,0.029556671,0.6599931],"study_design_scores_gemma":[0.000038094877,0.000039679093,0.00034768457,0.00006861755,0.00005496109,0.0001025677,0.00014963871,0.9200101,0.008311138,0.06259636,0.008258485,0.00002258261],"about_ca_topic_score_codex":0.0064456463,"about_ca_topic_score_gemma":0.012213602,"teacher_disagreement_score":0.0064456463,"about_ca_system_score_codex":0.0016317928,"about_ca_system_score_gemma":0.0014771845,"threshold_uncertainty_score":0.020451069},"labels":[],"label_agreement":null},{"id":"W3098980613","doi":"10.18653/v1/2020.emnlp-main.304","title":"Recurrent Interaction Network for Jointly Extracting Entities and Classifying Relations","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Beijing Advanced Innovation Center for Big Data and Brain Computing; Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China","keywords":"Computer science; Task (project management); Relation (database); Artificial intelligence; Machine learning; Multi-task learning; Relationship extraction; Task analysis; Joint (building); Data mining","score_opus":0.10479919264425545,"score_gpt":0.2967338189729715,"score_spread":0.19193462632871605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098980613","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028984036,0.0008565173,0.9647835,0.00036107007,0.00006997019,0.00007392588,0.00081736234,0.0015072015,0.00254644],"genre_scores_gemma":[0.7327457,0.0010873066,0.24947225,0.0002492353,0.00020737367,0.00034955936,0.006093619,0.00024671343,0.009548226],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991424,0.00024062903,0.000043619537,0.00030968682,0.00016872371,0.00009485733],"domain_scores_gemma":[0.9989329,0.00055293896,0.00017607433,0.00012548591,0.00016999825,0.000042492917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001195001,0.0012504803,0.00083583343,0.0020253882,0.00045545047,0.0008905936,0.0014153746,0.0009743162,0.0017641422],"category_scores_gemma":[0.003103645,0.00040484773,0.0012021816,0.0020934378,0.00047964277,0.0025706347,0.0010920219,0.0013402489,0.0008842723],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005643537,0.00027353456,0.0058606626,0.00026985258,0.00034106383,0.00050888804,0.00041774064,0.60411763,0.016893802,0.037886314,0.0126632275,0.3202029],"study_design_scores_gemma":[0.0000045309253,0.000019842788,0.00054354314,0.000005488394,0.000030382156,0.0000349946,0.000011130869,0.986602,0.001238496,0.0102151465,0.0012844283,0.000010028809],"about_ca_topic_score_codex":0.007856573,"about_ca_topic_score_gemma":0.009911601,"teacher_disagreement_score":0.007856573,"about_ca_system_score_codex":0.0009482612,"about_ca_system_score_gemma":0.00077496073,"threshold_uncertainty_score":0.015621722},"labels":[],"label_agreement":null},{"id":"W3099008231","doi":"10.18653/v1/2020.findings-emnlp.129","title":"exBERT: Extending Pre-trained Models with Domain-specific Vocabulary Under Constrained Training Resources","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Air Force Research Laboratory; MediaTek","keywords":"Computer science; Vocabulary; Benchmark (surveying); Artificial intelligence; Embedding; Domain (mathematical analysis); Context (archaeology); Machine learning; Computation; Training (meteorology); Training set; Natural language processing; Language model; Algorithm; Mathematics","score_opus":0.057954741990095024,"score_gpt":0.23438861524082547,"score_spread":0.17643387325073046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099008231","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061766032,0.003031231,0.8935305,0.001222719,0.0006618623,0.00034653745,0.0033931718,0.028570222,0.007477728],"genre_scores_gemma":[0.51637346,0.0019029017,0.42687732,0.0021667802,0.0005905197,0.00096353324,0.025749303,0.0027637526,0.022612423],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9992538,0.00021592833,0.000050701852,0.0002633539,0.00012526015,0.000090783215],"domain_scores_gemma":[0.9978544,0.0011111965,0.00007205544,0.0005513905,0.00031719572,0.0000936618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019099533,0.0025591238,0.0010248765,0.001337817,0.0005620045,0.0012811403,0.0027041167,0.0016441637,0.005444447],"category_scores_gemma":[0.0066230246,0.0011489635,0.0017513951,0.0011044992,0.00076977763,0.004164476,0.0027061442,0.003958076,0.004342335],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004489058,0.00037918388,0.00530784,0.00051072676,0.00046264438,0.00048337164,0.00037563473,0.38309208,0.012357337,0.007935319,0.048393577,0.5402534],"study_design_scores_gemma":[0.000047394995,0.00012351801,0.00067546393,0.00007008889,0.0000635672,0.00013682058,0.00006193091,0.9735346,0.004928887,0.010510318,0.009810834,0.0000365567],"about_ca_topic_score_codex":0.009831248,"about_ca_topic_score_gemma":0.022057416,"teacher_disagreement_score":0.009831248,"about_ca_system_score_codex":0.0009976545,"about_ca_system_score_gemma":0.0016315698,"threshold_uncertainty_score":0.019548059},"labels":[],"label_agreement":null},{"id":"W3099354896","doi":"10.18653/v1/2020.emnlp-main.681","title":"Deconstructing word embedding algorithms","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Word2vec; Word (group theory); Computer science; Word embedding; Natural language processing; Embedding; Feature (linguistics); Artificial intelligence; Resource (disambiguation); Quality (philosophy); Linguistics","score_opus":0.054864406494952514,"score_gpt":0.3044185296808561,"score_spread":0.2495541231859036,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099354896","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0109412065,0.0004898304,0.98657507,0.00033715044,0.000047667167,0.00004310721,0.00011812034,0.00045725668,0.00099067],"genre_scores_gemma":[0.31785566,0.0021011673,0.6683693,0.00042434063,0.00026949,0.0004053371,0.002260236,0.00077363814,0.007540901],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9972385,0.0013258277,0.00021465661,0.00064532174,0.00042147056,0.00015420598],"domain_scores_gemma":[0.9936231,0.0033337758,0.000311544,0.0015842677,0.0010064315,0.00014094028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035434074,0.0015366282,0.0009748297,0.0016227213,0.00051413657,0.0023954394,0.0013512765,0.0013555514,0.0022809869],"category_scores_gemma":[0.018277766,0.00067645893,0.0011187883,0.001428024,0.0016451136,0.0066305455,0.003500678,0.003462956,0.0017894479],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018077588,0.000119218304,0.0031568897,0.00032263654,0.00015139839,0.00011323913,0.0007595668,0.25222453,0.0077775037,0.26623175,0.005494153,0.4634685],"study_design_scores_gemma":[0.000013312626,0.00005991125,0.00033938797,0.000056428722,0.000018238368,0.000072831186,0.0001017072,0.7615878,0.0036699336,0.22827247,0.005784796,0.000023191718],"about_ca_topic_score_codex":0.0014957665,"about_ca_topic_score_gemma":0.0025685716,"teacher_disagreement_score":0.0035434074,"about_ca_system_score_codex":0.0007102387,"about_ca_system_score_gemma":0.0008506796,"threshold_uncertainty_score":0.018739522},"labels":[],"label_agreement":null},{"id":"W3099409970","doi":"10.1016/j.future.2020.11.012","title":"Entity-aware capsule network for multi-class classification of big data: A deep learning approach","year":2020,"lang":"en","type":"article","venue":"Future Generation Computer Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Research Center of the College of Computer and Information Sciences, King Saud University","keywords":"Computer science; Artificial intelligence; Convolutional neural network; Deep learning; Machine learning; Natural language processing; Pooling; Variety (cybernetics); Machine translation; Named-entity recognition; Task (project management)","score_opus":0.16835399261266557,"score_gpt":0.2749958446199877,"score_spread":0.10664185200732215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099409970","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.080292076,0.0017597995,0.91099,0.0014517385,0.00018910074,0.000104520026,0.0012257491,0.0015494309,0.0024375503],"genre_scores_gemma":[0.83201045,0.0014204456,0.15252279,0.0006188016,0.0003480434,0.00017962906,0.004695699,0.00017352361,0.008030572],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944025,0.00012552379,0.000030696254,0.00019181863,0.0000955391,0.00011605677],"domain_scores_gemma":[0.99867296,0.0005652263,0.00016421688,0.00025183847,0.00023346914,0.00011231956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015053542,0.0008749861,0.0011866078,0.0016363949,0.00072147994,0.0014994587,0.002775374,0.0016264918,0.0016077498],"category_scores_gemma":[0.0030068548,0.00050023955,0.0011074191,0.0025288947,0.00071754784,0.0036990265,0.0023840873,0.0020932278,0.0006106552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066762883,0.0007864887,0.014260723,0.00025161027,0.00045557105,0.0003747591,0.00041476716,0.3816838,0.005355861,0.03534769,0.026430624,0.5339705],"study_design_scores_gemma":[0.0000036486938,0.000015765447,0.00034921596,0.000007050952,0.00002246021,0.00001714945,0.000020526988,0.9914204,0.00043189558,0.007148911,0.00055776007,0.000005196806],"about_ca_topic_score_codex":0.011997934,"about_ca_topic_score_gemma":0.017112149,"teacher_disagreement_score":0.011997934,"about_ca_system_score_codex":0.0012169866,"about_ca_system_score_gemma":0.001333354,"threshold_uncertainty_score":0.023856223},"labels":[],"label_agreement":null},{"id":"W3099413717","doi":"10.18653/v1/2020.findings-emnlp.250","title":"Improving Word Embedding Factorization for Compression Using Distilled Nonlinear Neural Decomposition","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Embedding; Computer science; Word embedding; Word (group theory); Matrix decomposition; Artificial intelligence; Machine translation; Translation (biology); Language model; Speech recognition; Natural language processing; Mathematics","score_opus":0.04846563712461783,"score_gpt":0.31855323091445853,"score_spread":0.2700875937898407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099413717","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039194643,0.00044474637,0.9542438,0.00021702613,0.00012774793,0.00008267329,0.00030991057,0.0036846253,0.001694892],"genre_scores_gemma":[0.3490361,0.0005381319,0.6387652,0.0002924325,0.0001364293,0.00027536278,0.003138742,0.0007073675,0.0071102045],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99953353,0.000120274584,0.000034568522,0.00010981004,0.00014913669,0.00005261101],"domain_scores_gemma":[0.9986914,0.0005595107,0.000076817334,0.00030980483,0.0003119031,0.000050560164],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007679442,0.0014701415,0.000737073,0.00064950966,0.00031739107,0.00076904596,0.0007331843,0.0007939136,0.004380003],"category_scores_gemma":[0.0046918546,0.00030503207,0.0006819847,0.0008650372,0.0005399555,0.0024935922,0.0012050134,0.0016946646,0.002735092],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004351291,0.0003417042,0.0011808142,0.0002609464,0.00010427334,0.00022162162,0.00027262018,0.18314986,0.047698,0.013291961,0.011670514,0.7413725],"study_design_scores_gemma":[0.00003192975,0.00009179528,0.0002703043,0.000015466981,0.000021356671,0.000088461704,0.00005680613,0.96997327,0.02042335,0.006018216,0.0029926465,0.000016463322],"about_ca_topic_score_codex":0.0036160906,"about_ca_topic_score_gemma":0.0067411475,"teacher_disagreement_score":0.004380003,"about_ca_system_score_codex":0.00039987487,"about_ca_system_score_gemma":0.0008637967,"threshold_uncertainty_score":0.01465261},"labels":[],"label_agreement":null},{"id":"W3099471123","doi":"10.48550/arxiv.2011.10036","title":"On the Dynamics of Training Attention Models","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Discriminative model; Computer science; Classifier (UML); Artificial intelligence; Embedding; Machine learning; Stochastic gradient descent; Dynamics (music); Task (project management); Set (abstract data type); Block (permutation group theory); Simple (philosophy); Artificial neural network; Natural language processing; Psychology; Mathematics; Engineering","score_opus":0.18672170592516524,"score_gpt":0.18892304551873793,"score_spread":0.0022013395935726876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099471123","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22479203,0.002006042,0.7534718,0.00597207,0.00014650346,0.0000962107,0.0002687943,0.0007533889,0.012493177],"genre_scores_gemma":[0.95533735,0.00079150434,0.03699519,0.0005615658,0.00011123898,0.00017577808,0.0002754357,0.00024670057,0.0055051725],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991253,0.0003793338,0.00003328907,0.00021034124,0.00012672934,0.0001251168],"domain_scores_gemma":[0.98843646,0.009303534,0.0006385973,0.0005041118,0.0006934116,0.00042386257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029768846,0.0009338377,0.0012446749,0.0007103471,0.0007124922,0.0014742034,0.0013421375,0.001977806,0.0042813905],"category_scores_gemma":[0.036827113,0.00077456125,0.00046953632,0.0005319065,0.0018764337,0.003978,0.0024893912,0.0032864239,0.00054193026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027836204,0.00014026226,0.0063720536,0.00022426268,0.00008633573,0.00021144236,0.0007204336,0.71288866,0.0073723686,0.2119325,0.005375613,0.05439774],"study_design_scores_gemma":[0.000011466287,0.00002655137,0.0004066003,0.000019416118,0.0000077135555,0.000020578349,0.000024307974,0.9573498,0.00037015995,0.041416068,0.0003407047,0.0000066196108],"about_ca_topic_score_codex":0.0062024468,"about_ca_topic_score_gemma":0.0052218223,"teacher_disagreement_score":0.0062024468,"about_ca_system_score_codex":0.0017446921,"about_ca_system_score_gemma":0.0011069787,"threshold_uncertainty_score":0.015743434},"labels":[],"label_agreement":null},{"id":"W3099883129","doi":"","title":"Generating Extractive Summaries of Scientific Paradigms","year":2013,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval; Salient; Set (abstract data type); Citation; Dependency (UML); Data science; Cover (algebra); Parsing; Natural language processing; World Wide Web; Artificial intelligence","score_opus":0.026476075784331296,"score_gpt":0.2415757040701654,"score_spread":0.21509962828583412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099883129","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.084329166,0.0017570739,0.8693856,0.0013883408,0.00040852782,0.0007391813,0.016754173,0.017585406,0.007652449],"genre_scores_gemma":[0.18149121,0.0012486164,0.78020006,0.00013267307,0.0003228092,0.00064527715,0.030921122,0.0009935892,0.0040447167],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986399,0.00044963032,0.00018172529,0.00025247072,0.00042595805,0.000050393697],"domain_scores_gemma":[0.99094766,0.0049699503,0.0009712383,0.0009285483,0.0020246825,0.00015795462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017929228,0.0015648544,0.0006629245,0.0040422776,0.00071149407,0.0020456822,0.0010037137,0.0010094191,0.0035362074],"category_scores_gemma":[0.015929008,0.0005640599,0.0007996476,0.0032946796,0.000305946,0.0019649821,0.0012802033,0.00092430494,0.0024241335],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004533346,0.00029458367,0.0064775967,0.0022485608,0.00026687502,0.0014410015,0.0028645345,0.06642184,0.0671465,0.028353056,0.05170217,0.7723299],"study_design_scores_gemma":[0.00018104578,0.00057061645,0.0068799388,0.00032430416,0.0005019495,0.00095966045,0.0014070623,0.6669202,0.11245741,0.06616801,0.14347827,0.00015154057],"about_ca_topic_score_codex":0.0012300559,"about_ca_topic_score_gemma":0.0025459416,"teacher_disagreement_score":0.0040422776,"about_ca_system_score_codex":0.00066384696,"about_ca_system_score_gemma":0.0012055252,"threshold_uncertainty_score":0.011829734},"labels":[],"label_agreement":null},{"id":"W3100038904","doi":"10.18653/v1/2020.sustainlp-1.14","title":"A Little Bit Is Worse Than None: Ranking with Limited Training Data","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Training set; Ranking (information retrieval); Classifier (UML); Relevance (law); Labeled data; Artificial intelligence; Information retrieval; Machine learning; Domain (mathematical analysis); Simple (philosophy); Training (meteorology)","score_opus":0.1694170731781986,"score_gpt":0.2713215509351272,"score_spread":0.10190447775692862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3100038904","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5410271,0.009954116,0.4106563,0.009395684,0.0010889933,0.00035028358,0.0034450267,0.008144317,0.015938055],"genre_scores_gemma":[0.86036867,0.0007449573,0.1271468,0.0013564267,0.00041556187,0.000103228755,0.0045589367,0.00048752682,0.004817925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993292,0.0043095103,0.00026756595,0.0011073285,0.0007204583,0.00030313787],"domain_scores_gemma":[0.9772906,0.015642785,0.00066896086,0.004433276,0.0013422893,0.00062206166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010662103,0.0013618739,0.0027022595,0.0017373925,0.001518555,0.0024000201,0.0020798775,0.0021029036,0.0025310428],"category_scores_gemma":[0.037866864,0.000512941,0.00094485754,0.0013798586,0.0016464697,0.0064728633,0.0014693442,0.003332056,0.002169793],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030562659,0.0018337296,0.021670585,0.0011764612,0.0007130163,0.0003351274,0.00075508986,0.21775575,0.011887294,0.026573779,0.08271552,0.63152736],"study_design_scores_gemma":[0.00029863583,0.001778062,0.007041223,0.00016822194,0.00017340337,0.0005803488,0.0007537639,0.88148355,0.010161673,0.08585443,0.011519402,0.00018723216],"about_ca_topic_score_codex":0.0066496367,"about_ca_topic_score_gemma":0.012605461,"teacher_disagreement_score":0.010662103,"about_ca_system_score_codex":0.0011376516,"about_ca_system_score_gemma":0.0011891774,"threshold_uncertainty_score":0.056387305},"labels":[],"label_agreement":null},{"id":"W3100107515","doi":"10.18653/v1/2020.findings-emnlp.63","title":"Document Ranking with a Pretrained Sequence-to-Sequence Model","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":414,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Ranking (information retrieval); Sequence (biology); Encoder; Relevance (law); Artificial intelligence; Transformer; Task (project management); Information retrieval; Machine learning; Natural language processing; Data mining; Engineering","score_opus":0.07025758913945321,"score_gpt":0.27606782510455125,"score_spread":0.20581023596509804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3100107515","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0389448,0.0006051779,0.94900864,0.0007372547,0.00014748545,0.00023644247,0.0007849678,0.004130311,0.005404962],"genre_scores_gemma":[0.7078541,0.0006593933,0.25043646,0.00078615005,0.00035270958,0.0006784671,0.0028453714,0.000540605,0.035846747],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931395,0.00021988813,0.0000429348,0.00019907871,0.00014368999,0.000080555285],"domain_scores_gemma":[0.9982114,0.00095387106,0.0001320284,0.00028172397,0.00035174852,0.00006917467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016332672,0.0010342525,0.0011771971,0.0011720692,0.00043663647,0.0013702197,0.0019074052,0.0016481377,0.0063951164],"category_scores_gemma":[0.0042923475,0.00054601254,0.0009966518,0.0013169096,0.0006395476,0.0033149691,0.0006740945,0.0023127252,0.004700397],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035074016,0.00031217872,0.0013352908,0.00022943014,0.00010543942,0.00015105098,0.00009665953,0.6466268,0.00843034,0.018903572,0.011250964,0.31220755],"study_design_scores_gemma":[0.000013308905,0.00007242337,0.00015935824,0.0000069272974,0.000015222673,0.000047099285,0.0000077753075,0.9898121,0.0012434432,0.0077691074,0.000842837,0.000010457602],"about_ca_topic_score_codex":0.005607946,"about_ca_topic_score_gemma":0.011609267,"teacher_disagreement_score":0.0063951164,"about_ca_system_score_codex":0.0013160162,"about_ca_system_score_gemma":0.001346013,"threshold_uncertainty_score":0.021393836},"labels":[],"label_agreement":null},{"id":"W3100645984","doi":"10.18653/v1/2020.clinicalnlp-1.33","title":"Exploring Text Specific and Blackbox Fairness Algorithms in Multimodal Clinical NLP","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Vector Institute; University of Toronto","funders":"Vector Institute; Canadian Institute for Advanced Research","keywords":"Odds; Computer science; Intersection (aeronautics); Artificial intelligence; Word (group theory); Task (project management); Natural language processing; Modality (human–computer interaction); Machine learning; Algorithm; Logistic regression","score_opus":0.29510042896090366,"score_gpt":0.32190189009543474,"score_spread":0.02680146113453108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3100645984","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060636465,0.0005642084,0.9347637,0.0016080345,0.00009793024,0.00012218513,0.00015981072,0.00071228686,0.0013354513],"genre_scores_gemma":[0.66468793,0.00023796949,0.330445,0.0008195483,0.00035681523,0.0002206687,0.0004413595,0.00025599884,0.0025347755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9898645,0.006222795,0.00048004705,0.0019761862,0.0010173883,0.0004390526],"domain_scores_gemma":[0.9568721,0.03366833,0.002174775,0.00439377,0.0017792691,0.0011117299],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02590754,0.0010060075,0.001548784,0.0011816472,0.0014231964,0.0035004516,0.0023278988,0.0021744883,0.0024271125],"category_scores_gemma":[0.07715528,0.00042069712,0.0009195458,0.0010983889,0.0026263364,0.0059278,0.0055553094,0.0034348657,0.0006700314],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023754353,0.00087187934,0.021946806,0.000428938,0.00030729012,0.00029557693,0.0019737987,0.27597,0.008481211,0.16448437,0.0062178625,0.51664686],"study_design_scores_gemma":[0.00007682981,0.00017473908,0.0011255984,0.00003883255,0.00003321226,0.00008650417,0.00022800075,0.7919938,0.004562973,0.19975442,0.0018926684,0.000032426942],"about_ca_topic_score_codex":0.0017817841,"about_ca_topic_score_gemma":0.0016524562,"teacher_disagreement_score":0.02590754,"about_ca_system_score_codex":0.0017633992,"about_ca_system_score_gemma":0.0028225754,"threshold_uncertainty_score":0.13701385},"labels":[],"label_agreement":null},{"id":"W3100889051","doi":"10.48550/arxiv.2011.05723","title":"CalibreNet: Calibration Networks for Multilingual Sequence Labeling","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"People's Government of Jilin Province; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Benchmark (surveying); Sequence labeling; Task (project management); Sequence (biology); Named-entity recognition; Phrase; Natural language processing; Artificial intelligence; Obstacle; Boundary (topology); Resource (disambiguation); Mathematics","score_opus":0.1565945791399868,"score_gpt":0.22578983050627377,"score_spread":0.06919525136628696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3100889051","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01277017,0.0007105884,0.96122634,0.00038200803,0.00015661454,0.00017766158,0.0017244162,0.019103162,0.0037489722],"genre_scores_gemma":[0.24567409,0.00076245505,0.7093514,0.0011859855,0.000282191,0.0010538192,0.023795357,0.0029101144,0.014984586],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835145,0.0004714679,0.000064611406,0.0007695248,0.000207185,0.00013568734],"domain_scores_gemma":[0.9968311,0.0014456307,0.00019846483,0.00076928583,0.00061757053,0.00013791192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021335273,0.0023105126,0.0010810791,0.0020867167,0.001230547,0.0013730426,0.0033764248,0.002986632,0.007507518],"category_scores_gemma":[0.009836225,0.0010022469,0.0013831102,0.0020373738,0.0010259383,0.0049131718,0.0035175034,0.004204549,0.0057288636],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004603384,0.00033323342,0.0033194062,0.0003696908,0.00021255831,0.0003763465,0.00078389887,0.15792811,0.014953324,0.019321814,0.05421987,0.7477214],"study_design_scores_gemma":[0.00003121747,0.00007123704,0.00047371193,0.00005456971,0.0000325745,0.00010480238,0.00013348197,0.9489894,0.006105681,0.03208698,0.011881103,0.00003516699],"about_ca_topic_score_codex":0.0071223727,"about_ca_topic_score_gemma":0.0134843085,"teacher_disagreement_score":0.007507518,"about_ca_system_score_codex":0.0012623906,"about_ca_system_score_gemma":0.0018292399,"threshold_uncertainty_score":0.025115192},"labels":[],"label_agreement":null},{"id":"W3100949457","doi":"10.18653/v1/2020.findings-emnlp.127","title":"Filtering before Iteratively Referring for Knowledge-Grounded Response Selection in Retrieval-Based Chatbots","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Filter (signal processing); Context (archaeology); Knowledge extraction; Conversation; Artificial intelligence; Process (computing); Task (project management); Persona; Information retrieval; Selection (genetic algorithm); Knowledge base; Natural language processing; Machine learning; Human–computer interaction; Computer vision; Engineering","score_opus":0.07232172618019449,"score_gpt":0.3008768054510579,"score_spread":0.22855507927086338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3100949457","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.090466574,0.0017627091,0.87870103,0.00053120457,0.00016252015,0.00045071053,0.0006431454,0.022849571,0.0044324966],"genre_scores_gemma":[0.62251,0.0003518065,0.36119434,0.0006517458,0.00014920822,0.00040321908,0.0030331798,0.0008722908,0.0108342385],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967445,0.0012412304,0.00016073165,0.00096999714,0.0005181954,0.00036533896],"domain_scores_gemma":[0.99583435,0.002176986,0.00025772516,0.00080256234,0.00067194423,0.00025647884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034547332,0.0024256755,0.0019860202,0.0027831183,0.0014123026,0.0018392439,0.0032699984,0.002456227,0.004390241],"category_scores_gemma":[0.010680231,0.000695771,0.0015218715,0.0013349869,0.0011229995,0.0038639933,0.0027119492,0.0019310658,0.0036897403],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016731779,0.0009826008,0.008400897,0.0010679456,0.00035223793,0.0008815827,0.0041040606,0.038013775,0.05849234,0.010438929,0.025209563,0.85038286],"study_design_scores_gemma":[0.00014413142,0.0006249819,0.0043213195,0.000113899536,0.00030167296,0.0007119941,0.0015771794,0.9146433,0.04010943,0.022152264,0.015153895,0.00014598403],"about_ca_topic_score_codex":0.008265651,"about_ca_topic_score_gemma":0.01382388,"teacher_disagreement_score":0.008265651,"about_ca_system_score_codex":0.0010160224,"about_ca_system_score_gemma":0.001677312,"threshold_uncertainty_score":0.018270552},"labels":[],"label_agreement":null},{"id":"W3100985894","doi":"10.1109/micro50266.2020.00071","title":"GOBO: Quantizing Attention-Based NLP Models for Low Latency and Energy Efficient Inference","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":163,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Quantization (signal processing); Inference; Computation; Latency (audio); Language model; Parallel computing; Computer engineering; Computer hardware; Artificial intelligence; Algorithm","score_opus":0.11445244292019002,"score_gpt":0.2010061101867137,"score_spread":0.08655366726652367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3100985894","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017572105,0.0004563507,0.9743931,0.0003819868,0.00005954436,0.00007112404,0.00033346174,0.004760182,0.001972129],"genre_scores_gemma":[0.5540584,0.0005016627,0.43833154,0.00057039317,0.00009381857,0.00027648278,0.0012089033,0.0007525533,0.0042064013],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996779,0.000072163915,0.000016871123,0.000082940314,0.00011126491,0.000038815255],"domain_scores_gemma":[0.99931514,0.0003938483,0.000052149782,0.0001092018,0.00009777367,0.000031828793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000787173,0.00077370595,0.0007022862,0.0006487797,0.0003819125,0.0009373178,0.0018012304,0.00091738167,0.0047090584],"category_scores_gemma":[0.003968044,0.00044426217,0.0006022362,0.00066384376,0.0006738571,0.002464982,0.0016274336,0.0016186045,0.0010864995],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045985004,0.00015958452,0.001246104,0.00035004903,0.00013072633,0.00015423448,0.00027637457,0.55861,0.024826512,0.03394753,0.01230179,0.3675372],"study_design_scores_gemma":[0.000009884186,0.000012219165,0.00006574653,0.000006651376,0.0000059673857,0.000009789605,0.000009153959,0.987432,0.001909631,0.009809873,0.00072433427,0.0000046313885],"about_ca_topic_score_codex":0.010578421,"about_ca_topic_score_gemma":0.019946398,"teacher_disagreement_score":0.010578421,"about_ca_system_score_codex":0.0010602196,"about_ca_system_score_gemma":0.0010490513,"threshold_uncertainty_score":0.021033704},"labels":[],"label_agreement":null},{"id":"W3101035550","doi":"","title":"Measuring Systematic Generalization in Neural Proof Generation with Transformers","year":2020,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; McGill University","funders":"","keywords":"Mathematical proof; Computer science; Backward chaining; Generalization; Chaining; Artificial intelligence; Natural deduction; Automated theorem proving; Inference; Theoretical computer science; Programming language; Mathematics; Inference engine","score_opus":0.061497662312542785,"score_gpt":0.22974727517805896,"score_spread":0.16824961286551618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101035550","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.74801975,0.0008109179,0.24346194,0.00063226937,0.00006670438,0.00022637469,0.0007725187,0.003264302,0.0027452563],"genre_scores_gemma":[0.94889116,0.00020306098,0.049012113,0.0001475153,0.000013108119,0.00012628679,0.0010970585,0.000103366605,0.0004063666],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747473,0.00097149867,0.00024787357,0.0006880756,0.00045506484,0.00016263683],"domain_scores_gemma":[0.9604189,0.030643309,0.0021040065,0.004975502,0.00139482,0.00046349695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006694199,0.0011155056,0.0006130201,0.0010406238,0.000278985,0.001103836,0.0016156236,0.001363978,0.0016018722],"category_scores_gemma":[0.05228423,0.0006035756,0.0009938248,0.00068677263,0.001147759,0.0053343102,0.0015939804,0.00261706,0.0004656616],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007949703,0.0005756917,0.033208113,0.0006651753,0.00056920806,0.00018977755,0.00059469306,0.7490464,0.019704323,0.006530267,0.0018040432,0.18631741],"study_design_scores_gemma":[0.00003760072,0.00035858486,0.0034650392,0.000028830322,0.00006994122,0.00009852624,0.00006396481,0.97083336,0.011288314,0.01335307,0.00038044507,0.000022285252],"about_ca_topic_score_codex":0.0029241783,"about_ca_topic_score_gemma":0.0043085814,"teacher_disagreement_score":0.006694199,"about_ca_system_score_codex":0.001286137,"about_ca_system_score_gemma":0.0011397917,"threshold_uncertainty_score":0.035402715},"labels":[],"label_agreement":null},{"id":"W3101190870","doi":"10.18653/v1/2020.blackboxnlp-1.3","title":"Examining the rhetorical capacities of neural language models","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Canadian Institute for Advanced Research","keywords":"Rhetorical question; Computer science; ENCODE; Rhetoric; Set (abstract data type); Linguistics; Transformer; Rhetorical device; Artificial intelligence; Natural language processing; Philosophy; Programming language","score_opus":0.13039270456795185,"score_gpt":0.26006055282419666,"score_spread":0.1296678482562448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101190870","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8768371,0.00045850853,0.11071965,0.00117289,0.000034587963,0.000039912724,0.00023972451,0.00053852977,0.009959106],"genre_scores_gemma":[0.985886,0.00012462931,0.013114103,0.00004643296,0.000012994078,0.000045844623,0.00015986002,0.000031891483,0.0005781825],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988527,0.0007368699,0.00005319527,0.00013931294,0.00015285714,0.00006501058],"domain_scores_gemma":[0.97368956,0.023213279,0.0008960151,0.0010663309,0.00075890875,0.00037601465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005166506,0.0006387389,0.00048458943,0.0013275272,0.0004360868,0.0018700556,0.0007680005,0.0011476929,0.002449636],"category_scores_gemma":[0.041364674,0.00039037436,0.00040468146,0.0006954604,0.0012178737,0.005238059,0.0015017864,0.0014288722,0.00035363997],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006041179,0.00023452369,0.017987225,0.0006375045,0.00026818513,0.00022875474,0.001894961,0.7817342,0.026865935,0.06392603,0.0015922815,0.10402622],"study_design_scores_gemma":[0.00001213374,0.00008715945,0.0012627307,0.000017841157,0.000028159186,0.000023764413,0.00015449876,0.9793717,0.0038960406,0.014718887,0.0004137828,0.00001327034],"about_ca_topic_score_codex":0.0017313673,"about_ca_topic_score_gemma":0.0022281024,"teacher_disagreement_score":0.005166506,"about_ca_system_score_codex":0.0008630797,"about_ca_system_score_gemma":0.00058590644,"threshold_uncertainty_score":0.027323365},"labels":[],"label_agreement":null},{"id":"W3101384737","doi":"10.18653/v1/2020.findings-emnlp.281","title":"Towards Domain-Independent Text Structuring Trainable on Large Discourse Treebanks","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Structuring; Computer science; Pointer (user interface); Task (project management); Dependency (UML); Artificial intelligence; Metric (unit); Natural language processing; Set (abstract data type); Domain (mathematical analysis); Tree (set theory); Machine learning; Programming language","score_opus":0.023455984990626122,"score_gpt":0.26016114431718734,"score_spread":0.23670515932656122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101384737","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09729013,0.0010993576,0.8657681,0.0011227406,0.00014848194,0.0003921245,0.004378611,0.023952078,0.0058484497],"genre_scores_gemma":[0.30195054,0.00048702274,0.6743424,0.00041290806,0.00012715746,0.0006168774,0.016550176,0.0010280104,0.004484869],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99872595,0.0006008011,0.00007944907,0.00038859388,0.00013959607,0.00006552565],"domain_scores_gemma":[0.99272597,0.0055650664,0.00029272036,0.0006319013,0.0006140989,0.00017017948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029993197,0.0012752241,0.00082324335,0.0026163205,0.000790463,0.0015059155,0.0016304392,0.0018917127,0.004024601],"category_scores_gemma":[0.012111095,0.0006821056,0.00096566527,0.0023097566,0.0006529929,0.0058999215,0.002365189,0.0032067376,0.0024927882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004479373,0.00046922275,0.004835555,0.00093610224,0.00016070595,0.00024776108,0.0011963411,0.19687863,0.02665985,0.02175831,0.032617662,0.7137919],"study_design_scores_gemma":[0.000044425735,0.00006008464,0.00068477896,0.000060976898,0.000034978577,0.000033361033,0.0002129505,0.96667975,0.007959185,0.018749375,0.005465836,0.000014416476],"about_ca_topic_score_codex":0.0039038924,"about_ca_topic_score_gemma":0.010404465,"teacher_disagreement_score":0.004024601,"about_ca_system_score_codex":0.0013409498,"about_ca_system_score_gemma":0.0020103976,"threshold_uncertainty_score":0.015862107},"labels":[],"label_agreement":null},{"id":"W3101513017","doi":"10.18653/v1/2020.clinicalnlp-1.15","title":"MeDAL: Medical Abbreviation Disambiguation Dataset for Natural Language Understanding Pretraining","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University","funders":"","keywords":"Computer science; Natural language processing; Medal; Artificial intelligence; Domain (mathematical analysis); Natural language; Natural (archaeology)","score_opus":0.1476165516358032,"score_gpt":0.3630807248595844,"score_spread":0.21546417322378122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101513017","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043399468,0.004328857,0.03333676,0.0037459547,0.00081514765,0.0012965795,0.8861929,0.018025333,0.008859114],"genre_scores_gemma":[0.03383287,0.0005241134,0.045538675,0.000941171,0.00017270047,0.0009453306,0.91498816,0.00045603016,0.0026010396],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99799025,0.0005775734,0.000373057,0.00060082076,0.00034053664,0.000117829724],"domain_scores_gemma":[0.9953323,0.0024396523,0.00035717917,0.0008076669,0.0007083123,0.0003548859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022564384,0.0018793681,0.0009209864,0.0041246344,0.0010560729,0.0010173721,0.002201402,0.002973627,0.012351912],"category_scores_gemma":[0.011184453,0.00041914117,0.0012856245,0.0021681983,0.0006687702,0.0013163246,0.0019623712,0.0022573017,0.010751927],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070074026,0.0004089542,0.0080253845,0.0033064312,0.00023165769,0.0012096221,0.00046879743,0.0050827963,0.010616916,0.0028272492,0.87181896,0.09530247],"study_design_scores_gemma":[0.0012408858,0.0005270647,0.031519603,0.00065342494,0.00029467393,0.005010434,0.0008216132,0.044768084,0.028756335,0.010846059,0.87531835,0.00024351756],"about_ca_topic_score_codex":0.006441092,"about_ca_topic_score_gemma":0.018824201,"teacher_disagreement_score":0.012351912,"about_ca_system_score_codex":0.001278889,"about_ca_system_score_gemma":0.0027111305,"threshold_uncertainty_score":0.041321218},"labels":[],"label_agreement":null},{"id":"W3101878262","doi":"10.18653/v1/2020.emnlp-main.325","title":"Supervised Seeded Iterated Learning for Interactive Language Learning","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Compute Canada","keywords":"Computer science; Artificial intelligence; Iterated function; Language model; Natural language processing; Strengths and weaknesses; Task (project management); Word (group theory); Seeding; Supervised learning; Language acquisition; Dynamics (music); Train; Natural language; Machine learning; Linguistics; Artificial neural network; Psychology; Mathematics","score_opus":0.02719451474617501,"score_gpt":0.2676619374156743,"score_spread":0.24046742266949928,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101878262","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033097178,0.00019207566,0.9625273,0.00021136827,0.000029281882,0.0000748284,0.000044464734,0.0017560353,0.002067541],"genre_scores_gemma":[0.7607934,0.00009699803,0.2343452,0.00022448518,0.000057752164,0.0002852655,0.00018189353,0.000280435,0.0037345788],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99857736,0.00067191175,0.00005926166,0.00031940147,0.00024894474,0.00012317408],"domain_scores_gemma":[0.99514705,0.0035511593,0.00026919928,0.0005053112,0.0003302532,0.00019709452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028093208,0.001066982,0.00086044346,0.0005560354,0.00056663936,0.0009003982,0.0024810112,0.0015104035,0.0032116582],"category_scores_gemma":[0.009555941,0.000519923,0.00069640746,0.0003797703,0.0015384188,0.0022555722,0.0022826153,0.002192202,0.00078277907],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033676106,0.0003466393,0.0024562785,0.00020585417,0.0001241556,0.00039746522,0.0007854809,0.7618779,0.012918745,0.045141693,0.003424029,0.171985],"study_design_scores_gemma":[0.000010265996,0.000030232111,0.0000413155,0.0000033845397,0.000003364102,0.000014642372,0.000008639354,0.99216014,0.00089921046,0.006575475,0.00024914442,0.0000042310867],"about_ca_topic_score_codex":0.0029096953,"about_ca_topic_score_gemma":0.0043642297,"teacher_disagreement_score":0.0032116582,"about_ca_system_score_codex":0.0011508316,"about_ca_system_score_gemma":0.0011420633,"threshold_uncertainty_score":0.014857233},"labels":[],"label_agreement":null},{"id":"W3102645206","doi":"10.18653/v1/2020.emnlp-main.506","title":"Factual Error Correction for Abstractive Summarization Models","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":124,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Automatic summarization; Computer science; Consistency (knowledge bases); Artificial intelligence; Heuristic; Machine learning; Natural language processing","score_opus":0.08399768128692162,"score_gpt":0.2738264388879877,"score_spread":0.18982875760106607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3102645206","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038173176,0.0016805772,0.9372132,0.00081509963,0.0002797342,0.0002657457,0.0010390182,0.017864991,0.0026684676],"genre_scores_gemma":[0.55214554,0.0008639775,0.43204793,0.0004337544,0.00033470124,0.00036716068,0.004788333,0.0009385832,0.008080024],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998415,0.00045301585,0.00019367423,0.00052747683,0.0003249489,0.00008591939],"domain_scores_gemma":[0.99348736,0.0028448903,0.000883603,0.001213144,0.0014307164,0.00014034203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029888675,0.0013931467,0.000864869,0.0013521435,0.0004678881,0.0015934615,0.0021679045,0.0013400331,0.0027047456],"category_scores_gemma":[0.014499954,0.00035439056,0.00077310856,0.0007526793,0.0005709507,0.0023384332,0.0010261935,0.0018407686,0.0016464179],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043048835,0.000193539,0.0034524747,0.0006726644,0.00026779078,0.0002594552,0.00057439436,0.2674629,0.02318385,0.0076089157,0.019813662,0.67607987],"study_design_scores_gemma":[0.000023054306,0.000113641036,0.0007105112,0.000044775643,0.000063264815,0.00006835562,0.00005831069,0.971826,0.016290005,0.0052744146,0.005504992,0.000022680959],"about_ca_topic_score_codex":0.0042945407,"about_ca_topic_score_gemma":0.0072915866,"teacher_disagreement_score":0.0042945407,"about_ca_system_score_codex":0.0011855192,"about_ca_system_score_gemma":0.0012253427,"threshold_uncertainty_score":0.015806794},"labels":[],"label_agreement":null},{"id":"W3103026670","doi":"10.18653/v1/2020.sdp-1.4","title":"A Smart System to Generate and Validate Question Answer Pairs for COVID-19 Literature","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; Thomson Reuters (Canada); Queen's University","funders":"","keywords":"Computer science; Annotation; Coronavirus disease 2019 (COVID-19); Domain (mathematical analysis); Transformer; Task (project management); Subject-matter expert; Information retrieval; Data science; Artificial intelligence; World Wide Web","score_opus":0.03858974180858433,"score_gpt":0.27716645475049223,"score_spread":0.2385767129419079,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3103026670","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036189824,0.0007826192,0.4979724,0.0015109965,0.00054739905,0.0028358514,0.042793162,0.40714332,0.010224484],"genre_scores_gemma":[0.14844294,0.00032297213,0.7312281,0.0008439401,0.00020646497,0.0025840458,0.102004565,0.0053896504,0.008977222],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99554235,0.001330466,0.0005564576,0.001679742,0.00075060694,0.00014033036],"domain_scores_gemma":[0.98507476,0.008033361,0.00096522167,0.00271545,0.002499403,0.00071175874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008150267,0.0019095042,0.0011937346,0.009562546,0.0012075423,0.0025833468,0.002050726,0.0024849714,0.016888117],"category_scores_gemma":[0.029763965,0.00092904054,0.0015818818,0.0027524305,0.00056149584,0.006101384,0.0044823484,0.0017665126,0.017559476],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014103138,0.0010801388,0.020195195,0.0028514958,0.00037618034,0.0012657908,0.003738177,0.0073327865,0.053501688,0.012718738,0.20714545,0.68838406],"study_design_scores_gemma":[0.0004200128,0.0010728682,0.014858943,0.00042872428,0.00031645477,0.0019211405,0.0017603931,0.4797606,0.09799321,0.032039426,0.3690946,0.0003336635],"about_ca_topic_score_codex":0.0020943596,"about_ca_topic_score_gemma":0.0035898157,"teacher_disagreement_score":0.016888117,"about_ca_system_score_codex":0.0014297277,"about_ca_system_score_gemma":0.0022384387,"threshold_uncertainty_score":0.05649638},"labels":[],"label_agreement":null},{"id":"W3103356223","doi":"10.18653/v1/2020.emnlp-main.101","title":"Learning VAE-LDA Models with Rounded Reparameterization Trick","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Artificial intelligence; Generative grammar; Benchmark (surveying); Topic model; Dirichlet distribution; Generative model; Machine learning; Mathematics","score_opus":0.044944351810645704,"score_gpt":0.2241349230300437,"score_spread":0.179190571219398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3103356223","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008427133,0.0006016762,0.9886602,0.0003112059,0.000054560267,0.000056903085,0.00014235836,0.0007184747,0.0010274277],"genre_scores_gemma":[0.38613862,0.0013818573,0.59884584,0.000893851,0.00043512083,0.0011519893,0.00242228,0.0006284705,0.008101927],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99764794,0.0013394816,0.00011014651,0.00049537787,0.00026099276,0.0001460666],"domain_scores_gemma":[0.9959942,0.0026707833,0.00020932942,0.00071023207,0.00029710625,0.000118416036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039777863,0.001322132,0.002066175,0.0013937929,0.0008109632,0.0021861326,0.0028651527,0.0020648914,0.003954111],"category_scores_gemma":[0.013244155,0.0010501732,0.0022027944,0.0016141022,0.0013000024,0.0044459305,0.003230331,0.004902567,0.0032614633],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002849474,0.0002372937,0.0019975638,0.00029652056,0.0003067832,0.00019668917,0.00056457875,0.47669658,0.0053964686,0.11657607,0.014532126,0.38291436],"study_design_scores_gemma":[0.000017256945,0.000024596,0.00010023045,0.000016665168,0.000012255178,0.000047153553,0.000023337936,0.94873226,0.0006553841,0.048687574,0.0016652101,0.000018028057],"about_ca_topic_score_codex":0.0016830242,"about_ca_topic_score_gemma":0.003005057,"teacher_disagreement_score":0.0039777863,"about_ca_system_score_codex":0.00082391355,"about_ca_system_score_gemma":0.0009325422,"threshold_uncertainty_score":0.021036804},"labels":[],"label_agreement":null},{"id":"W3103410567","doi":"10.48550/arxiv.2011.06188","title":"Evaluating Curriculum Learning Strategies in Neural Combinatorial Optimization","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Computer science; Curriculum; Machine learning; Artificial intelligence; Inefficiency; Task (project management); Field (mathematics); Artificial neural network; Range (aeronautics); Engineering; Mathematics; Systems engineering; Psychology","score_opus":0.11831965325499277,"score_gpt":0.24173151748570462,"score_spread":0.12341186423071185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3103410567","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3864032,0.0022239937,0.5947461,0.0016484674,0.00013079403,0.0003827934,0.00010059236,0.0005536998,0.01381033],"genre_scores_gemma":[0.86160386,0.0007492331,0.13450591,0.0002683469,0.00007579582,0.00046877205,0.0001811174,0.000111653026,0.0020353268],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998346,0.0011455684,0.000058712223,0.00017307377,0.00017675044,0.000099979974],"domain_scores_gemma":[0.97803736,0.019216362,0.00073232077,0.0007613399,0.0007812155,0.0004713968],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064609763,0.0010447949,0.0009200286,0.0009365331,0.0004935389,0.0013370852,0.0015969275,0.001807154,0.0020415005],"category_scores_gemma":[0.037854556,0.0004713397,0.00042450044,0.00079734554,0.0014758589,0.0024032176,0.0017583601,0.0020793285,0.00030863687],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018493863,0.0003482028,0.0020892955,0.00022787164,0.00007315619,0.000025633799,0.00012510447,0.8926775,0.0006104638,0.033058923,0.0011055914,0.0694734],"study_design_scores_gemma":[0.0000318431,0.00010293097,0.00015911019,0.000019286743,0.000007512826,0.0000037835273,0.000018436738,0.98565567,0.0003289864,0.013414481,0.00025234558,0.0000054543925],"about_ca_topic_score_codex":0.0020708228,"about_ca_topic_score_gemma":0.0027318385,"teacher_disagreement_score":0.0064609763,"about_ca_system_score_codex":0.001605446,"about_ca_system_score_gemma":0.0013815055,"threshold_uncertainty_score":0.034169316},"labels":[],"label_agreement":null},{"id":"W3103720397","doi":"10.2196/23357","title":"Using Character-Level and Entity-Level Representations to Enhance Bidirectional Encoder Representation From Transformers-Based Clinical Semantic Textual Similarity Model: ClinicalSTS Modeling Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Harbin Institute of Technology","keywords":"Computer science; Natural language processing; Snippet; Artificial intelligence; Information retrieval; Encoder; Semantic similarity; Transformer; Sentence; Language model","score_opus":0.2662994500200416,"score_gpt":0.45955414980796283,"score_spread":0.19325469978792126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3103720397","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21362764,0.0006083919,0.77651656,0.00091484253,0.00007592985,0.00024761542,0.001478331,0.0024503863,0.004080373],"genre_scores_gemma":[0.8666537,0.00031485545,0.12480205,0.00019526259,0.0000480592,0.00022008951,0.0027505602,0.000126189,0.0048892666],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99947387,0.00015777283,0.000036737318,0.00017182485,0.00010440505,0.000055533248],"domain_scores_gemma":[0.9983865,0.0009902347,0.0001521681,0.00016337557,0.00024507975,0.0000625708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008278855,0.0006913557,0.00040177602,0.0011102225,0.00022915729,0.00069846265,0.0012875452,0.00084896287,0.0023079012],"category_scores_gemma":[0.004327206,0.00023081593,0.0010318039,0.00088301935,0.00038776305,0.001685282,0.00068642653,0.0010911906,0.0005744732],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005300832,0.00045806024,0.00980642,0.00018306087,0.00014562158,0.00043803075,0.0002955546,0.7070616,0.009515078,0.02010962,0.0043494417,0.24710743],"study_design_scores_gemma":[0.0000058734954,0.000028933084,0.00031539728,0.0000038416492,0.000012122695,0.00003235775,0.000010505324,0.995948,0.0011556486,0.0021157402,0.00036702224,0.000004491981],"about_ca_topic_score_codex":0.012790964,"about_ca_topic_score_gemma":0.013755906,"teacher_disagreement_score":0.012790964,"about_ca_system_score_codex":0.0013560796,"about_ca_system_score_gemma":0.0012734993,"threshold_uncertainty_score":0.025433064},"labels":[],"label_agreement":null},{"id":"W3103849289","doi":"","title":"Complex Question Answering: Unsupervised Learning Approaches and Experiments","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Lethbridge","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Automatic summarization; Set (abstract data type); Information retrieval; Weighting; Cosine similarity; Tree kernel; Cluster analysis; Kernel method; Support vector machine","score_opus":0.21775882767795485,"score_gpt":0.27658948023646857,"score_spread":0.05883065255851372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3103849289","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7577524,0.002661062,0.20775041,0.00092282053,0.00031634799,0.003373244,0.0063660587,0.0058370573,0.015020623],"genre_scores_gemma":[0.77124524,0.00073442963,0.20793831,0.00054237514,0.00016471627,0.0022323593,0.012747731,0.00048603537,0.003908739],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9921292,0.0043648914,0.00061875687,0.00128853,0.001248515,0.00035010662],"domain_scores_gemma":[0.96883595,0.023342526,0.00095421064,0.004118569,0.0021396258,0.00060913246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007941198,0.0014115685,0.001385071,0.0013459977,0.0011500699,0.0012147363,0.0023780626,0.0023410365,0.0039718603],"category_scores_gemma":[0.024716498,0.0004702893,0.00093832205,0.002090566,0.0015219947,0.0026065253,0.0019466681,0.0023122127,0.0017249419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049066353,0.01720871,0.02385549,0.0031884634,0.0010966115,0.0008949928,0.002496938,0.38881275,0.018816162,0.008956588,0.03627273,0.49349385],"study_design_scores_gemma":[0.0006850559,0.0014941265,0.011087151,0.00011296649,0.00012757479,0.0004974423,0.00067317946,0.94555354,0.0159433,0.0151484385,0.008533508,0.00014365072],"about_ca_topic_score_codex":0.00677223,"about_ca_topic_score_gemma":0.005786841,"teacher_disagreement_score":0.007941198,"about_ca_system_score_codex":0.0012977837,"about_ca_system_score_gemma":0.0010714128,"threshold_uncertainty_score":0.04199761},"labels":[],"label_agreement":null},{"id":"W3104616748","doi":"","title":"Unsupervised Text Generation by Learning from Search","year":2020,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Paraphrase; Computer science; Artificial intelligence; Generative grammar; Bootstrapping (finance); Unsupervised learning; Machine learning; Baseline (sea); Beam search; Natural language processing; Simulated annealing; Search algorithm; Algorithm; Mathematics","score_opus":0.06081246761507966,"score_gpt":0.25057610381942186,"score_spread":0.1897636362043422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3104616748","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020950504,0.0003995904,0.9743188,0.00024388148,0.000033843015,0.000106014646,0.0002194635,0.0019686029,0.0017592807],"genre_scores_gemma":[0.55747044,0.00036139222,0.43215895,0.00042424447,0.00011764025,0.00048035252,0.0019282884,0.00070197886,0.006356771],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993819,0.00023475161,0.000034400564,0.00019476945,0.00011350612,0.000040629482],"domain_scores_gemma":[0.9982101,0.0010820208,0.00013868476,0.00031087457,0.0002017809,0.000056552144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00084638706,0.0005131038,0.0007582895,0.001069299,0.00038756756,0.0006920046,0.0016162683,0.00091032794,0.0026379945],"category_scores_gemma":[0.004300535,0.00034437896,0.0007870057,0.0010599046,0.0007401941,0.0018013383,0.00095098955,0.0009499861,0.0009771563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021905774,0.0002705323,0.0024935335,0.0003641485,0.00012953348,0.00022374507,0.00039408286,0.5307276,0.01315667,0.06605622,0.012716651,0.3732482],"study_design_scores_gemma":[0.000020050413,0.000025312449,0.0001323994,0.000007853248,0.000009542006,0.00003478275,0.000013967735,0.9728282,0.0021057068,0.023217907,0.0015974734,0.0000069247585],"about_ca_topic_score_codex":0.0015981254,"about_ca_topic_score_gemma":0.004321685,"teacher_disagreement_score":0.0026379945,"about_ca_system_score_codex":0.00070049375,"about_ca_system_score_gemma":0.00089776434,"threshold_uncertainty_score":0.0088249445},"labels":[],"label_agreement":null},{"id":"W3104738015","doi":"10.18653/v1/2020.sustainlp-1.11","title":"Early Exiting BERT for Efficient Document Ranking","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Vector Institute","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Speedup; Inference; Ranking (information retrieval); Latency (audio); Computation; Language model; Code (set theory); Artificial intelligence; Parallel computing; Algorithm; Programming language; Set (abstract data type); Telecommunications","score_opus":0.03341623580851892,"score_gpt":0.2571902084438234,"score_spread":0.22377397263530446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3104738015","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022664469,0.0011112073,0.90442216,0.0003437098,0.00022276859,0.00022289353,0.0013033549,0.0661926,0.0035167397],"genre_scores_gemma":[0.3408755,0.0005193771,0.6324785,0.000514139,0.0002933759,0.0003453282,0.0062775095,0.003852718,0.014843525],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99876803,0.00035938143,0.000072295705,0.00024183336,0.0003680541,0.00019042857],"domain_scores_gemma":[0.99713886,0.0013118614,0.0001272621,0.00083883654,0.00042360753,0.00015954362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017296153,0.0017015032,0.0012548628,0.001236173,0.0008935836,0.0018844294,0.0021558641,0.0012268481,0.012497129],"category_scores_gemma":[0.00813763,0.000862199,0.00092431816,0.0012045852,0.0005944972,0.0036502783,0.001952584,0.0026507187,0.010658488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013565252,0.00032486185,0.0036763866,0.000459322,0.00013882545,0.00029494896,0.00035605024,0.10428921,0.023993328,0.017381899,0.050341554,0.7973871],"study_design_scores_gemma":[0.00007296451,0.00014700477,0.00052283995,0.000021380534,0.000031012725,0.00016143761,0.000057968366,0.95935863,0.012122403,0.015647266,0.011810995,0.0000460922],"about_ca_topic_score_codex":0.009610074,"about_ca_topic_score_gemma":0.026382843,"teacher_disagreement_score":0.012497129,"about_ca_system_score_codex":0.0009775879,"about_ca_system_score_gemma":0.0022347642,"threshold_uncertainty_score":0.041807055},"labels":[],"label_agreement":null},{"id":"W3104938695","doi":"10.18653/v1/2020.sdp-1.19","title":"Cydex: Neural Search Infrastructure for the Scholarly Literature","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Open source; Search engine; Generality; Ranking (information retrieval); Digital library; Domain (mathematical analysis); World Wide Web; Software; Programming language","score_opus":0.03427892096320958,"score_gpt":0.2669885034713679,"score_spread":0.2327095825081583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3104938695","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010035403,0.0030597975,0.07652004,0.0013567142,0.0004092564,0.0006382349,0.6137369,0.24505481,0.049188886],"genre_scores_gemma":[0.033362504,0.0015213278,0.11842974,0.00050090806,0.0001329035,0.0009992128,0.8193135,0.007715536,0.018024474],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99911135,0.00014458684,0.000111512374,0.00025746328,0.00026994222,0.00010526071],"domain_scores_gemma":[0.9981906,0.0005977131,0.00015368545,0.00044031892,0.0004189529,0.00019865591],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0011578398,0.0018107487,0.00095011835,0.0107120145,0.00092410244,0.003228575,0.0021118915,0.0011494795,0.054675713],"category_scores_gemma":[0.007728042,0.0006132489,0.0012301198,0.009570534,0.00038115654,0.00486756,0.004163108,0.0014089819,0.055823516],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036650282,0.0001093138,0.002273101,0.0020930553,0.00013908079,0.00020118774,0.00040460058,0.003001426,0.004347921,0.0142247025,0.8467992,0.12603983],"study_design_scores_gemma":[0.00017602262,0.00010763403,0.0045710444,0.00038797446,0.00013024722,0.00035493376,0.0005255543,0.07389486,0.009886968,0.035624675,0.8742053,0.00013488684],"about_ca_topic_score_codex":0.016977208,"about_ca_topic_score_gemma":0.039564047,"teacher_disagreement_score":0.99677145,"about_ca_system_score_codex":0.0020266648,"about_ca_system_score_gemma":0.0030563662,"threshold_uncertainty_score":0.18290848},"labels":[],"label_agreement":null},{"id":"W3105160568","doi":"10.18653/v1/2020.findings-emnlp.259","title":"Inferring symmetry in natural language","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Predicate (mathematical logic); Computer science; Natural language processing; Inference; Artificial intelligence; Symmetry (geometry); Context (archaeology); Natural language; Sentence; Linguistics; Mathematics; Programming language; Philosophy","score_opus":0.018308665181922135,"score_gpt":0.25028394417563515,"score_spread":0.231975278993713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3105160568","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.089568175,0.0006114325,0.9002828,0.00093670445,0.000058997342,0.00020437322,0.0023920394,0.0013624874,0.004582998],"genre_scores_gemma":[0.7163134,0.00032497183,0.27598587,0.00026318076,0.00012748345,0.00036821183,0.0054415124,0.00020120524,0.0009741242],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966162,0.0015228823,0.00025767626,0.0009562702,0.000506756,0.0001403094],"domain_scores_gemma":[0.98847187,0.006859528,0.0017989669,0.0019064604,0.0007765705,0.00018664304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004278426,0.0006520722,0.0005133944,0.0046050237,0.0011170114,0.0018902556,0.0014163373,0.00091033016,0.0036081688],"category_scores_gemma":[0.024514874,0.0005215615,0.0015403624,0.0026760192,0.002013679,0.006862254,0.0026977821,0.0016451175,0.000910583],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032980795,0.00033769722,0.04897651,0.0010015625,0.00042726254,0.00094098586,0.0067224894,0.043677554,0.022416623,0.50289524,0.010686009,0.36158827],"study_design_scores_gemma":[0.000043886535,0.00011123383,0.009651007,0.000120991004,0.00010605466,0.00059880986,0.0011404941,0.27979663,0.007988449,0.6809271,0.01944016,0.00007520624],"about_ca_topic_score_codex":0.0023813997,"about_ca_topic_score_gemma":0.0021390577,"teacher_disagreement_score":0.0046050237,"about_ca_system_score_codex":0.0010661768,"about_ca_system_score_gemma":0.0014959566,"threshold_uncertainty_score":0.022626698},"labels":[],"label_agreement":null},{"id":"W3105424285","doi":"10.18653/v1/2020.emnlp-main.393","title":"Utility is in the Eye of the User: A Critique of NLP Leaderboards","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":137,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Framing (construction); Inference; Artificial intelligence; Natural language processing; Machine learning","score_opus":0.05896975066937212,"score_gpt":0.2934283674829232,"score_spread":0.2344586168135511,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3105424285","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011047062,0.0059210057,0.46926966,0.424247,0.0014212596,0.00016938319,0.0004132931,0.0014781975,0.086033195],"genre_scores_gemma":[0.6969238,0.004092801,0.21514569,0.057629853,0.004354284,0.0013932134,0.0003125782,0.0028391047,0.017308692],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.92569536,0.050973885,0.0025366154,0.007610002,0.012044235,0.001139917],"domain_scores_gemma":[0.69304943,0.25570804,0.006505097,0.026525045,0.015459715,0.0027526706],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.083459504,0.0017569762,0.0021881224,0.0034801252,0.004928741,0.0133881625,0.007542246,0.011502569,0.0070416196],"category_scores_gemma":[0.26342022,0.0013777877,0.002209244,0.0037754818,0.03945647,0.024729433,0.009041471,0.021687461,0.002598386],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029814992,0.000018589437,0.00028940002,0.0001220815,0.000022032178,0.00005922073,0.0021319108,0.0027857553,0.00005866815,0.9714697,0.01132719,0.011685674],"study_design_scores_gemma":[0.000039961575,0.000016643648,0.00010556805,0.00016661998,0.000010407077,0.00003782277,0.000411294,0.012189773,0.00018268447,0.9467349,0.040075388,0.000029015937],"about_ca_topic_score_codex":0.008319877,"about_ca_topic_score_gemma":0.0049322736,"teacher_disagreement_score":0.9165405,"about_ca_system_score_codex":0.0095345555,"about_ca_system_score_gemma":0.0073039117,"threshold_uncertainty_score":0.44138134},"labels":[],"label_agreement":null},{"id":"W3105971000","doi":"","title":"Experiments with Three Approaches to Recognizing Lexical Entailment","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Similarity (geometry); Word (group theory); Context (archaeology); Feature vector; Feature (linguistics); Set (abstract data type); Word embedding; Semantic similarity; Concatenation (mathematics); Relation (database); Mathematics; Data mining; Image (mathematics)","score_opus":0.23513063562285139,"score_gpt":0.26334298577569026,"score_spread":0.028212350152838878,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3105971000","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77925694,0.012238801,0.11786547,0.0049052266,0.0018554569,0.008436397,0.019281417,0.026781823,0.029378472],"genre_scores_gemma":[0.41750082,0.002134171,0.5103018,0.0023759904,0.00036892493,0.0038727156,0.055864386,0.00075246295,0.006828696],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9808926,0.0065533547,0.0035611764,0.0048570847,0.0033542837,0.000781488],"domain_scores_gemma":[0.9413884,0.043104317,0.0016774003,0.0073746024,0.0047430643,0.0017121896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014873058,0.0044959355,0.0026902424,0.004550725,0.0028877023,0.003734799,0.009630423,0.0076974747,0.00634801],"category_scores_gemma":[0.057204075,0.0013603139,0.0037936394,0.0053108977,0.002377498,0.0109808305,0.004729098,0.0077076876,0.0029669239],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011991343,0.016909923,0.021447176,0.009255373,0.0030582077,0.0015825698,0.0024025652,0.0789644,0.026974503,0.010040754,0.08529519,0.73207796],"study_design_scores_gemma":[0.0059803836,0.007072409,0.022781776,0.00067734346,0.0017404289,0.0021694947,0.0046834257,0.7923493,0.07383114,0.030412603,0.057677157,0.00062456273],"about_ca_topic_score_codex":0.015224475,"about_ca_topic_score_gemma":0.024032623,"teacher_disagreement_score":0.015224475,"about_ca_system_score_codex":0.0036520166,"about_ca_system_score_gemma":0.0033361858,"threshold_uncertainty_score":0.07865721},"labels":[],"label_agreement":null},{"id":"W3108142956","doi":"10.3233/faia200860","title":"Sentence Embeddings and High-Speed Similarity Search for Fast Computer Assisted Annotation of Legal Documents","year":2020,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Research Unit on Children's Psychosocial Maladjustment","funders":"","keywords":"Annotation; Computer science; Sentence; Natural language processing; Artificial intelligence; Similarity (geometry); Meaning (existential); Process (computing); Information retrieval; Interface (matter); Programming language; Image (mathematics)","score_opus":0.053547645764866245,"score_gpt":0.2979148153250844,"score_spread":0.24436716956021814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3108142956","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022324352,0.00087011396,0.95978135,0.0003161651,0.00025084338,0.00019543008,0.0010298092,0.012062014,0.0031698735],"genre_scores_gemma":[0.16857783,0.000386707,0.81966233,0.00014425744,0.00017254493,0.000223347,0.005344509,0.0008003323,0.004688113],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985006,0.0005825837,0.00013370563,0.0003645095,0.00035227754,0.0000662815],"domain_scores_gemma":[0.9973821,0.00113445,0.00023177885,0.0005331841,0.0006221112,0.000096417294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014780556,0.00081324915,0.00068349816,0.002216819,0.0006437697,0.0011723193,0.0010202806,0.0010592852,0.0065200925],"category_scores_gemma":[0.0070179803,0.0004532598,0.00061825983,0.0023910974,0.0004348707,0.0039439886,0.0018614002,0.0012207336,0.005075087],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044490714,0.00027963438,0.0016191262,0.0005819115,0.00011123342,0.00036292616,0.001024762,0.010693379,0.04611864,0.029986084,0.049805652,0.85897183],"study_design_scores_gemma":[0.000087471715,0.00021567916,0.0025983923,0.000072227616,0.00007150981,0.0006722923,0.0005413854,0.85584426,0.03493674,0.055103447,0.04976516,0.00009145058],"about_ca_topic_score_codex":0.0021446387,"about_ca_topic_score_gemma":0.0036758396,"teacher_disagreement_score":0.0065200925,"about_ca_system_score_codex":0.0005655678,"about_ca_system_score_gemma":0.0008255607,"threshold_uncertainty_score":0.021811903},"labels":[],"label_agreement":null},{"id":"W3109501393","doi":"10.2196/21750","title":"Family History Information Extraction With Neural Attention and an Enhanced Relation-Side Scheme: Algorithm Development and Validation","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Ministry of Science and Technology, Taiwan","keywords":"Computer science; Artificial neural network; Artificial intelligence; Relationship extraction; Information extraction; Machine learning; Scheme (mathematics); Task (project management); Inference; Relation (database); Sentence; Algorithm; Data mining","score_opus":0.026119777972762416,"score_gpt":0.2653891485497481,"score_spread":0.2392693705769857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109501393","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32717496,0.0025812855,0.65306693,0.0011732887,0.00023898546,0.0012286043,0.0011597974,0.0087486105,0.00462752],"genre_scores_gemma":[0.57803106,0.0005368612,0.41393116,0.00044171698,0.000070848786,0.00073174527,0.0024602774,0.00012042099,0.0036758864],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99920756,0.00021020134,0.00007978282,0.00028258582,0.000121170284,0.00009870898],"domain_scores_gemma":[0.9973143,0.0015310112,0.00013017727,0.00025282617,0.0006970304,0.00007464339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030438907,0.0013174904,0.0010584878,0.001303078,0.00065650704,0.0008872696,0.0021521202,0.0019763394,0.002699847],"category_scores_gemma":[0.0068379045,0.0005393616,0.0008174112,0.001009738,0.00046725522,0.0017283355,0.0013937345,0.002132852,0.0006609342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087133364,0.0005924767,0.0072145537,0.00025577043,0.00022740396,0.00021022985,0.00013572606,0.3140067,0.00743898,0.0015656046,0.004545049,0.6629362],"study_design_scores_gemma":[0.000037079837,0.00006058225,0.0005458188,0.000010527155,0.000030149691,0.000027775885,0.000019732846,0.9960019,0.002372614,0.00060280116,0.0002848417,0.000006241959],"about_ca_topic_score_codex":0.025929295,"about_ca_topic_score_gemma":0.021884037,"teacher_disagreement_score":0.025929295,"about_ca_system_score_codex":0.0019893385,"about_ca_system_score_gemma":0.0024656667,"threshold_uncertainty_score":0.051556766},"labels":[],"label_agreement":null},{"id":"W3109558947","doi":"10.1609/aaai.v35i14.17527","title":"DialogBERT: Discourse-Aware Response Generation via Learning to Recover and Rank Utterances","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University","keywords":"Computer science; Utterance; Security token; Transformer; Natural language processing; Coherence (philosophical gambling strategy); Artificial intelligence; Conversation; Speech recognition; Language model; Psychology; Communication","score_opus":0.07664952376012893,"score_gpt":0.3068271529254254,"score_spread":0.23017762916529647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109558947","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020331366,0.00045258683,0.93991417,0.00029655933,0.0001668291,0.0003732147,0.0006733037,0.034518346,0.0032735746],"genre_scores_gemma":[0.43814814,0.00025592136,0.5354759,0.00054887857,0.00015562598,0.00092337426,0.0031380297,0.0016489357,0.01970513],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989774,0.00048252984,0.0000239457,0.00029180292,0.00014570117,0.000078742436],"domain_scores_gemma":[0.9988217,0.00066267385,0.00006080066,0.00019869242,0.00015636397,0.000099698475],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017812304,0.0016961964,0.000710213,0.00074740656,0.0004543393,0.0007317335,0.0020352295,0.0013232726,0.0071984855],"category_scores_gemma":[0.0039185295,0.00052212627,0.00068706984,0.0003500481,0.00054200925,0.0014922731,0.0019870128,0.0018409865,0.0038962378],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011073264,0.00054527284,0.0014848452,0.00045283904,0.00014218717,0.00025964945,0.0009884704,0.07738575,0.04024076,0.007228626,0.034671325,0.8354929],"study_design_scores_gemma":[0.00007806466,0.00023504233,0.00051608693,0.00002579487,0.000033893462,0.00012448491,0.00016471974,0.96389383,0.017562324,0.0077187014,0.009601683,0.00004546156],"about_ca_topic_score_codex":0.0030550617,"about_ca_topic_score_gemma":0.0058852714,"teacher_disagreement_score":0.0071984855,"about_ca_system_score_codex":0.00062286924,"about_ca_system_score_gemma":0.00091303507,"threshold_uncertainty_score":0.02408135},"labels":[],"label_agreement":null},{"id":"W3109752747","doi":"10.18653/v1/2020.coling-main.230","title":"Improving Conversational Question Answering Systems after Deployment using Feedback-Weighted Learning","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"CHIST-ERA; Ministerio de Educación, Cultura y Deporte; Nvidia; Samsung Advanced Institute of Technology; Canadian Institute for Advanced Research; Eusko Jaurlaritza; Agencia Estatal de Investigación; Samsung; Agence Nationale de la Recherche","keywords":"Computer science; Software deployment; Exploit; Domain (mathematical analysis); Binary number; Binary classification; Recommender system; Question answering; Matching (statistics); Artificial intelligence; Machine learning; Information retrieval; Human–computer interaction; Software engineering; Support vector machine","score_opus":0.029392632998733945,"score_gpt":0.2517073156640895,"score_spread":0.22231468266535553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109752747","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25727922,0.0013357966,0.70278585,0.0012022433,0.0002584084,0.0007599893,0.0005335044,0.032803655,0.003041346],"genre_scores_gemma":[0.7130038,0.00016174838,0.28066587,0.0004336687,0.00013696191,0.00037774586,0.0018287854,0.00078830204,0.0026032247],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99126196,0.005308904,0.00034593022,0.0016601967,0.0009951371,0.0004279081],"domain_scores_gemma":[0.9769613,0.013632949,0.00075373636,0.0033533485,0.004380698,0.00091784314],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010768855,0.0017432758,0.0012906959,0.0009622675,0.0008009521,0.0013623552,0.0023270804,0.0018112503,0.0019167855],"category_scores_gemma":[0.03789665,0.00071589905,0.00073549326,0.0005299175,0.0007231034,0.0043397727,0.0027493057,0.0027922182,0.0020824804],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002040595,0.0021457798,0.012207934,0.000920771,0.00039384997,0.0003087242,0.0024777737,0.09896893,0.09707049,0.0018997478,0.016093977,0.76547146],"study_design_scores_gemma":[0.00011985503,0.0005241397,0.0020459804,0.000025249827,0.000060223803,0.00009532557,0.00025343007,0.9640558,0.02654944,0.0028426354,0.0033789969,0.00004901164],"about_ca_topic_score_codex":0.004476488,"about_ca_topic_score_gemma":0.0062054666,"teacher_disagreement_score":0.010768855,"about_ca_system_score_codex":0.001006965,"about_ca_system_score_gemma":0.0013961297,"threshold_uncertainty_score":0.05695188},"labels":[],"label_agreement":null},{"id":"W3109919947","doi":"10.2196/22508","title":"Identification of Semantically Similar Sentences in Clinical Notes: Iterative Intermediate Training Using Multi-Task Learning","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Task (project management); Identification (biology); Natural language processing; Training (meteorology); Artificial intelligence; Engineering","score_opus":0.10714068958502107,"score_gpt":0.37892976319237814,"score_spread":0.27178907360735705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109919947","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26928553,0.001610868,0.70840365,0.0013026596,0.00043195547,0.0010886744,0.001498991,0.012767307,0.0036103935],"genre_scores_gemma":[0.7374456,0.0002738328,0.2505531,0.0009698808,0.00013631492,0.00091062044,0.0069056773,0.0004032712,0.0024017135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956863,0.0017022067,0.00033053994,0.0015585343,0.0003977315,0.0003246389],"domain_scores_gemma":[0.9853495,0.010342477,0.0005567943,0.0013725815,0.0018120288,0.00056669774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060936618,0.0029540816,0.0011521055,0.0018649235,0.0013317113,0.0016815964,0.0033169172,0.0026364066,0.0023549013],"category_scores_gemma":[0.019309396,0.0008664436,0.002328674,0.0012240718,0.0012738985,0.0035571202,0.004421824,0.0051831733,0.0019514691],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001484182,0.002031368,0.018088067,0.00083141844,0.00045205507,0.0012126996,0.0020749248,0.13493197,0.025572341,0.0021087711,0.017006757,0.7942055],"study_design_scores_gemma":[0.0001207199,0.0006441438,0.0021064118,0.00007899251,0.00011073614,0.0002859879,0.00052715867,0.97287273,0.014013794,0.0069884714,0.0021843275,0.00006659339],"about_ca_topic_score_codex":0.005943497,"about_ca_topic_score_gemma":0.008689156,"teacher_disagreement_score":0.0060936618,"about_ca_system_score_codex":0.0014682345,"about_ca_system_score_gemma":0.0025638323,"threshold_uncertainty_score":0.03222674},"labels":[],"label_agreement":null},{"id":"W3110578295","doi":"10.1186/s13640-020-00539-x","title":"Named entity recognition for Chinese judgment documents based on BiLSTM and CRF","year":2020,"lang":"en","type":"article","venue":"EURASIP Journal on Image and Video Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Division of Graduate Education; Natural Science Foundation of Guangxi Province; National Natural Science Foundation of China","keywords":"Computer science; Conditional random field; Named-entity recognition; Sentence; Task (project management); Artificial intelligence; Natural language processing; Paragraph; Domain (mathematical analysis); Field (mathematics); Hidden Markov model; Viterbi algorithm; Speech recognition; Pattern recognition (psychology)","score_opus":0.03360616821490845,"score_gpt":0.29459733110444974,"score_spread":0.2609911628895413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110578295","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28091457,0.0020816817,0.6718595,0.000667932,0.0005996475,0.00038701057,0.011052567,0.024360532,0.00807648],"genre_scores_gemma":[0.6795921,0.00067379576,0.28314415,0.00014851979,0.00014932934,0.0002533737,0.028703809,0.00040195946,0.006933094],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99899524,0.00020610524,0.00009432353,0.00044367166,0.00015785541,0.00010275727],"domain_scores_gemma":[0.99869317,0.0004166061,0.00010414219,0.00027529863,0.00044422413,0.000066631685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010773891,0.00093808246,0.0007594154,0.0019193845,0.0008004263,0.00077667186,0.0014116963,0.00092529855,0.003711447],"category_scores_gemma":[0.0029976976,0.00029803487,0.000712521,0.0023451361,0.00035784105,0.0024954479,0.0007726439,0.001118478,0.002544783],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075713143,0.0003246428,0.007535305,0.00063896953,0.0001376214,0.0013034809,0.00043517764,0.054370087,0.054446705,0.008067008,0.04860295,0.8233809],"study_design_scores_gemma":[0.00004458165,0.00010230479,0.005306154,0.00002608226,0.00006109264,0.00028425397,0.00016050487,0.94523406,0.032999814,0.0045995996,0.011113879,0.00006773913],"about_ca_topic_score_codex":0.015937053,"about_ca_topic_score_gemma":0.014389658,"teacher_disagreement_score":0.015937053,"about_ca_system_score_codex":0.0008384934,"about_ca_system_score_gemma":0.0011748391,"threshold_uncertainty_score":0.03168857},"labels":[],"label_agreement":null},{"id":"W3110664450","doi":"10.1609/aaai.v35i15.17601","title":"MASKER: Masked Keyword Regularization for Reliable Text Classification","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Defense Acquisition Program Administration; Agency for Defense Development","keywords":"Computer science; Regularization (linguistics); Artificial intelligence; Inference; Natural language processing; Context (archaeology); Generalization; Domain (mathematical analysis); Code (set theory); Pattern recognition (psychology); Speech recognition; Machine learning; Mathematics; Set (abstract data type)","score_opus":0.10858455132910795,"score_gpt":0.2992483494078932,"score_spread":0.19066379807878525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110664450","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033733457,0.0006561287,0.9444998,0.00066014536,0.0002405967,0.00015286476,0.0007749643,0.01750641,0.0017756298],"genre_scores_gemma":[0.3464859,0.0003622357,0.63830674,0.0008092491,0.00030763826,0.0005019471,0.0030953262,0.0021538998,0.007977019],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99858826,0.00039613634,0.00009729808,0.00043516548,0.00034744595,0.00013562405],"domain_scores_gemma":[0.9975714,0.001067624,0.00024741414,0.0006111818,0.0003799964,0.00012232432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002424701,0.0016997133,0.0010976718,0.0010942959,0.0006611468,0.0011613204,0.0022730366,0.0019581935,0.0041497382],"category_scores_gemma":[0.009747388,0.0005732201,0.0010943656,0.00083332596,0.0008050945,0.0030225369,0.0019233257,0.0026339216,0.004216004],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012120511,0.00052135414,0.004077036,0.00046135177,0.00022324796,0.00048614602,0.00067497394,0.13445495,0.09920325,0.015822034,0.049389344,0.69347435],"study_design_scores_gemma":[0.000042434225,0.00008153687,0.0005100917,0.000022621209,0.000018956363,0.00011516705,0.000044428,0.96620643,0.018345742,0.010356054,0.004224127,0.000032490676],"about_ca_topic_score_codex":0.0030331467,"about_ca_topic_score_gemma":0.005786602,"teacher_disagreement_score":0.0041497382,"about_ca_system_score_codex":0.0008247124,"about_ca_system_score_gemma":0.0014556202,"threshold_uncertainty_score":0.01388222},"labels":[],"label_agreement":null},{"id":"W3110846353","doi":"10.1609/aaai.v35i16.17680","title":"Reinforced Multi-Teacher Selection for Knowledge Distillation","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":121,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Bottleneck; Distillation; Computer science; Artificial intelligence; Inference; Machine learning; Selection (genetic algorithm); Learning cycle; Natural language processing; Mathematics education; Mathematics","score_opus":0.13650346173500516,"score_gpt":0.3305311318721846,"score_spread":0.19402767013717942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110846353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04144417,0.00058570376,0.9503056,0.00040724373,0.00008085468,0.00012369428,0.00013909097,0.004950334,0.001963233],"genre_scores_gemma":[0.6683879,0.00027227588,0.32399166,0.0005194155,0.0001299748,0.00038534874,0.0006924685,0.00059415586,0.00502689],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99910957,0.00034928808,0.000039789986,0.00023832251,0.0001478297,0.00011517447],"domain_scores_gemma":[0.9980246,0.0011893607,0.00011687443,0.0003154099,0.0002441794,0.00010969101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015297367,0.0015097182,0.0013457381,0.00074015395,0.0007453372,0.0008011486,0.002488753,0.0015710306,0.0029447987],"category_scores_gemma":[0.0056873737,0.0007635451,0.00080057245,0.00087435317,0.00096026325,0.0025008519,0.0020130666,0.0025026489,0.0012442328],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062676944,0.00036275288,0.0018790275,0.00021372347,0.000121869314,0.0002624163,0.00034201427,0.45255697,0.014576033,0.011970977,0.007217062,0.5098704],"study_design_scores_gemma":[0.000026978829,0.00004439475,0.000092757815,0.0000063191214,0.000011760656,0.00002520927,0.000015067966,0.9910313,0.003254833,0.0047097616,0.0007734829,0.00000821286],"about_ca_topic_score_codex":0.0033638137,"about_ca_topic_score_gemma":0.0076831346,"teacher_disagreement_score":0.0033638137,"about_ca_system_score_codex":0.0008102839,"about_ca_system_score_gemma":0.001584682,"threshold_uncertainty_score":0.0098513365},"labels":[],"label_agreement":null},{"id":"W3110879614","doi":"","title":"On the Systematicity of Probing Contextualized Word Representations: The Case of Hypernymy in BERT","year":2020,"lang":"en","type":"article","venue":"Joint Conference on Lexical and Computational Semantics","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; McGill University","funders":"","keywords":"Computer science; Competence (human resources); Noun; ENCODE; Consistency (knowledge bases); Natural language processing; Artificial intelligence; Cognitive science; Linguistics; Psychology; Philosophy","score_opus":0.11018733260050881,"score_gpt":0.2956192324816643,"score_spread":0.18543189988115547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110879614","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54668784,0.0012928959,0.4154182,0.0068724314,0.00010883418,0.00022338406,0.0005232123,0.001267049,0.02760614],"genre_scores_gemma":[0.9582821,0.00018934035,0.039852675,0.0004695962,0.00003519712,0.00007630565,0.0002788072,0.00017583679,0.00064014446],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9858233,0.007331509,0.0010199553,0.0032521028,0.0019701463,0.0006031002],"domain_scores_gemma":[0.823885,0.12831561,0.009478915,0.031867817,0.0051656323,0.0012869979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017707776,0.00080478314,0.00085541554,0.002360029,0.0014570472,0.004374564,0.0015014353,0.0030148302,0.003672172],"category_scores_gemma":[0.1575872,0.0010632905,0.0009038812,0.0016959469,0.00972139,0.020551251,0.007441875,0.0041855955,0.0005945196],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001164211,0.0001962216,0.09247816,0.0013678835,0.00033068008,0.0018955634,0.09534524,0.008657454,0.09258183,0.33160102,0.003748655,0.37063307],"study_design_scores_gemma":[0.00009497158,0.00045817404,0.063598186,0.0006802791,0.00023786147,0.0036396696,0.017553642,0.058778554,0.05367806,0.77811897,0.022794235,0.00036748493],"about_ca_topic_score_codex":0.0027443236,"about_ca_topic_score_gemma":0.0022217298,"teacher_disagreement_score":0.017707776,"about_ca_system_score_codex":0.00084435433,"about_ca_system_score_gemma":0.0010138009,"threshold_uncertainty_score":0.09364879},"labels":[],"label_agreement":null},{"id":"W3110986327","doi":"10.1609/aaai.v35i16.17713","title":"CARE: Commonsense-Aware Emotional Response Generation with Latent Concepts","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Research (Canada)","funders":"Nanyang Technological University; National Research Foundation Singapore; National Research Foundation","keywords":"Rationality; Construct (python library); Computer science; Field (mathematics); Cognitive psychology; Commonsense reasoning; Neglect; Focus (optics); Commonsense knowledge; Artificial intelligence; Cognitive science; Psychology; Epistemology; Knowledge base","score_opus":0.10740496312420665,"score_gpt":0.3075754435097127,"score_spread":0.20017048038550606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110986327","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030953364,0.00050747226,0.9628081,0.0007353191,0.00010045953,0.00024844464,0.00017721552,0.0021850572,0.0022846316],"genre_scores_gemma":[0.7478228,0.00023440474,0.24505255,0.0006837361,0.00011470206,0.0006226506,0.0005136988,0.00020130014,0.0047540898],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980033,0.001089452,0.00006574924,0.00048830314,0.0002335577,0.00011960229],"domain_scores_gemma":[0.9960381,0.0027305863,0.0002704719,0.00047316667,0.00027851074,0.00020912812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002839928,0.0010960209,0.00075119286,0.0006647184,0.00049521413,0.0011451413,0.0017994643,0.0014050277,0.0029271313],"category_scores_gemma":[0.010598033,0.00047129023,0.0012537767,0.00042204678,0.0011382619,0.0024070388,0.0022785638,0.0024627866,0.0008203326],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011974162,0.0010909695,0.008624764,0.00055196404,0.00038566903,0.0004736735,0.0030651442,0.35457864,0.030654414,0.059188675,0.015243148,0.52494556],"study_design_scores_gemma":[0.00004935995,0.00009346866,0.0005471979,0.000016286564,0.000033558186,0.000063805106,0.00007047705,0.96926033,0.0024652383,0.025625149,0.0017477258,0.00002737973],"about_ca_topic_score_codex":0.0021735902,"about_ca_topic_score_gemma":0.0031526273,"teacher_disagreement_score":0.0029271313,"about_ca_system_score_codex":0.00092618313,"about_ca_system_score_gemma":0.00069577363,"threshold_uncertainty_score":0.015019119},"labels":[],"label_agreement":null},{"id":"W3111446732","doi":"10.1609/aaai.v35i14.17549","title":"Unsupervised Learning of Discourse Structures using a Tree Autoencoder","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Huawei Technologies","keywords":"Computer science; Autoencoder; Parsing; Artificial intelligence; Natural language processing; Tree (set theory); Set (abstract data type); Tree structure; Task (project management); Process (computing); Parse tree; Annotation; Deep learning; Binary tree","score_opus":0.10659551351510793,"score_gpt":0.323993007079681,"score_spread":0.21739749356457305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111446732","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03166223,0.00030240128,0.96444625,0.00025736078,0.000054149867,0.000054937806,0.00021756337,0.0016592091,0.0013459175],"genre_scores_gemma":[0.46746126,0.0004888614,0.52341026,0.00026644816,0.00011285042,0.00026632292,0.0018424573,0.0002627765,0.0058886902],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995763,0.00014634419,0.000019158118,0.00016417552,0.000054226777,0.00003982885],"domain_scores_gemma":[0.99891114,0.0007778173,0.0000628704,0.00007593055,0.00014304659,0.000029226132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009272683,0.00067007204,0.00058303645,0.0008375887,0.00036932915,0.00064742914,0.0008196868,0.00084250944,0.0012637356],"category_scores_gemma":[0.0022567904,0.00047941832,0.0008473639,0.00063016504,0.00051237654,0.0012614712,0.00084953866,0.0018840913,0.0007438606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021087365,0.00020521093,0.0018869514,0.00019303741,0.00016794013,0.00022016287,0.00052345585,0.34149712,0.03280178,0.017998481,0.004896228,0.59939873],"study_design_scores_gemma":[0.0000058822357,0.000017690476,0.00023934766,0.000011037231,0.000012233389,0.000013150909,0.000018147544,0.9918223,0.0030311777,0.0041151503,0.00070843747,0.0000054386433],"about_ca_topic_score_codex":0.003409156,"about_ca_topic_score_gemma":0.008003305,"teacher_disagreement_score":0.003409156,"about_ca_system_score_codex":0.0006354104,"about_ca_system_score_gemma":0.0009271981,"threshold_uncertainty_score":0.006778598},"labels":[],"label_agreement":null},{"id":"W3111733452","doi":"10.1109/smc42975.2020.9282902","title":"A Graph Based Approach to Automate Essay Evaluation","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brandon University; Lakehead University","funders":"","keywords":"Computer science; Consistency (knowledge bases); Natural language processing; Artificial intelligence; Machine learning; Reliability (semiconductor); Similarity (geometry); Process (computing); Graph; Semantic similarity; Random forest; Grammar; Computational linguistics; Power (physics); Programming language; Theoretical computer science","score_opus":0.07298428983306357,"score_gpt":0.2803077348305304,"score_spread":0.20732344499746683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111733452","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013533091,0.0002172306,0.9692002,0.00014228027,0.000075882956,0.00039944155,0.00095277687,0.011947513,0.003531636],"genre_scores_gemma":[0.22347005,0.00021051576,0.7657037,0.0000864988,0.00008035953,0.0003606331,0.0034958776,0.0007071946,0.0058851894],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99731714,0.0009556103,0.00023218585,0.00052566663,0.0008483332,0.00012112456],"domain_scores_gemma":[0.99359655,0.0024437155,0.00044536736,0.00072822126,0.002619737,0.00016635434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020878897,0.0013007425,0.00083203503,0.006777402,0.00081823097,0.001633872,0.001342755,0.00092574157,0.004488517],"category_scores_gemma":[0.009248209,0.0003442439,0.00089660863,0.003246389,0.00035554726,0.0016613527,0.0009996693,0.0009450057,0.0028581684],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025730432,0.00028439733,0.005794001,0.0003279718,0.00014804475,0.00022001854,0.00025554944,0.034015693,0.024753435,0.008519773,0.013004465,0.9124193],"study_design_scores_gemma":[0.000034452027,0.00020976027,0.0053188573,0.00004023301,0.00007037893,0.0002606091,0.00015975718,0.93836385,0.019071473,0.01921322,0.017201243,0.00005602552],"about_ca_topic_score_codex":0.007217198,"about_ca_topic_score_gemma":0.013755835,"teacher_disagreement_score":0.007217198,"about_ca_system_score_codex":0.00081223116,"about_ca_system_score_gemma":0.0010843668,"threshold_uncertainty_score":0.015015602},"labels":[],"label_agreement":null},{"id":"W3112090146","doi":"","title":"Knowledge graphs meet moral values","year":2020,"lang":"en","type":"article","venue":"MADOC (University of Mannheim)","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; Deutsche Forschungsgemeinschaft","keywords":"Operationalization; Morality; Lexicon; Epistemology; Relevance (law); Computer science; Core (optical fiber); Foundation (evidence); Sociology; Artificial intelligence; Political science; Philosophy; Law","score_opus":0.050939116035885926,"score_gpt":0.2145841937861747,"score_spread":0.16364507775028875,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3112090146","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08781122,0.0015200317,0.8516331,0.0048054163,0.00020142997,0.0003233314,0.003979879,0.0011028977,0.048622765],"genre_scores_gemma":[0.76321787,0.0016756605,0.2140105,0.0008417442,0.00031963742,0.0005093306,0.008273699,0.0003359904,0.010815531],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9965713,0.001461115,0.00025944118,0.00092626805,0.0005619761,0.0002199186],"domain_scores_gemma":[0.9868216,0.008869573,0.0008294268,0.0017567726,0.0012013499,0.0005213046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017075326,0.0006979689,0.0007346573,0.0041553103,0.0019615449,0.003902083,0.001006461,0.0021448745,0.009070652],"category_scores_gemma":[0.019587127,0.0005154153,0.0011386789,0.004218357,0.002403363,0.009878592,0.0028808755,0.0017280154,0.001460785],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011053213,0.000095111696,0.0043795984,0.00065459235,0.00015134714,0.0007111607,0.0031351964,0.029669588,0.0020995655,0.8107832,0.014875406,0.13333473],"study_design_scores_gemma":[0.000017008435,0.000012060614,0.00090831635,0.00007644477,0.000047074165,0.00024748011,0.0007317955,0.04304537,0.0006479004,0.92714393,0.027107306,0.000015271931],"about_ca_topic_score_codex":0.0048962343,"about_ca_topic_score_gemma":0.00427581,"teacher_disagreement_score":0.009070652,"about_ca_system_score_codex":0.0016315449,"about_ca_system_score_gemma":0.0012166991,"threshold_uncertainty_score":0.030344367},"labels":[],"label_agreement":null},{"id":"W3112648112","doi":"10.1186/s12911-020-01330-8","title":"CERC: an interactive content extraction, recognition, and construction tool for clinical and biomedical text","year":2020,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canada Excellence Research Chairs, Government of Canada; Georgia Institute of Technology; National Science Foundation","keywords":"Automatic summarization; Computer science; Relevance (law); Natural language processing; Artificial intelligence; Ranking (information retrieval); Visualization; Information retrieval; Test set; Multi-document summarization; Vocabulary; Set (abstract data type); Random forest","score_opus":0.21349124466008843,"score_gpt":0.4023797459099472,"score_spread":0.18888850124985876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3112648112","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015060974,0.0009002992,0.57178664,0.00088123494,0.00018704208,0.0014906671,0.019886086,0.38654774,0.0032593152],"genre_scores_gemma":[0.04738561,0.00042030375,0.9088655,0.0004453848,0.00022701957,0.0018355829,0.030976506,0.006456084,0.003387983],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9979353,0.00056163437,0.0002458665,0.00047599446,0.0006868765,0.00009436481],"domain_scores_gemma":[0.98351914,0.010747799,0.0017220452,0.0010751283,0.002466549,0.00046919452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048487196,0.0025751463,0.0010059036,0.008529502,0.0006929051,0.0016156827,0.0017651056,0.0013877479,0.016280973],"category_scores_gemma":[0.017685024,0.0006229405,0.0012237483,0.0027882035,0.00055611803,0.0026948778,0.0022158024,0.0010631654,0.008522968],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008730385,0.0002162926,0.003326488,0.0021129034,0.00026736065,0.0010134733,0.001183353,0.005121855,0.03895284,0.0022177126,0.2147142,0.7300005],"study_design_scores_gemma":[0.0011349461,0.0011834268,0.025769267,0.0011179287,0.0005514173,0.004798783,0.0012589316,0.49186593,0.14960766,0.02123614,0.30076918,0.00070638955],"about_ca_topic_score_codex":0.0018822396,"about_ca_topic_score_gemma":0.0029623571,"teacher_disagreement_score":0.016280973,"about_ca_system_score_codex":0.00073357567,"about_ca_system_score_gemma":0.0017114956,"threshold_uncertainty_score":0.054465294},"labels":[],"label_agreement":null},{"id":"W3112776819","doi":"10.1609/aaai.v35i14.17485","title":"Segatron: Segment-Aware Transformer for Language Modeling and Understanding","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Peng Cheng Laboratory","keywords":"Computer science; Transformer; Language model; Security token; Perplexity; Sentence; Artificial intelligence; Natural language processing; Speech recognition; Voltage; Engineering","score_opus":0.14032639263289431,"score_gpt":0.30855227845219885,"score_spread":0.16822588581930453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3112776819","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014765628,0.0005723407,0.7880742,0.00035271567,0.00023762102,0.00028651961,0.0127467625,0.17846559,0.0044985833],"genre_scores_gemma":[0.21348387,0.00083473546,0.6855058,0.0005805014,0.00009818526,0.0008086819,0.073481426,0.010361533,0.01484518],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944156,0.00012495341,0.000040837916,0.00025067193,0.000085601125,0.000056343477],"domain_scores_gemma":[0.9991672,0.00035330682,0.000042689502,0.00025451047,0.0001324684,0.000049865124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009042917,0.0024046665,0.00070253684,0.0015245466,0.0005266982,0.0014423679,0.0025811412,0.0010963117,0.015980463],"category_scores_gemma":[0.0032291997,0.0009990685,0.0021034682,0.0011061149,0.0005420442,0.0042814254,0.0018771251,0.0027593276,0.011124589],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071316963,0.00030164432,0.0031992604,0.0010486825,0.00033534341,0.00045181366,0.0006482136,0.07660989,0.028828707,0.021049235,0.16648339,0.7003306],"study_design_scores_gemma":[0.000082140235,0.0001256624,0.0005332865,0.000043756852,0.000073508025,0.0002662509,0.00014374465,0.9187769,0.021260653,0.021739809,0.03689483,0.000059406568],"about_ca_topic_score_codex":0.011175122,"about_ca_topic_score_gemma":0.026260972,"teacher_disagreement_score":0.015980463,"about_ca_system_score_codex":0.0013130564,"about_ca_system_score_gemma":0.001986592,"threshold_uncertainty_score":0.053459942},"labels":[],"label_agreement":null},{"id":"W3113945743","doi":"10.18653/v1/2020.semeval-1.76","title":"TR at SemEval-2020 Task 4: Exploring the Limits of Language-model-based Common Sense Validation","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Commonsense knowledge; SemEval; Concreteness; Language model; Natural language processing; Commonsense reasoning; Task (project management); Artificial intelligence; Embedding; Knowledge graph; Language understanding; Common sense; Domain knowledge; Cognitive psychology; Psychology","score_opus":0.09066099485163504,"score_gpt":0.28276391145479113,"score_spread":0.1921029166031561,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3113945743","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34471092,0.009099332,0.4330403,0.02100757,0.009367794,0.0031874038,0.036570907,0.08741638,0.05559934],"genre_scores_gemma":[0.64823055,0.000507142,0.23815109,0.0049304357,0.00095534103,0.0019775606,0.081657685,0.010062304,0.0135277845],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94638586,0.03791442,0.0020078968,0.007152085,0.005073199,0.0014665507],"domain_scores_gemma":[0.8050822,0.13578168,0.0031022048,0.03236105,0.019010521,0.004662351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0500135,0.0046358155,0.0026178963,0.0026823222,0.0030586729,0.0076133986,0.0073084454,0.007604038,0.014528429],"category_scores_gemma":[0.17353283,0.0011284712,0.0026248067,0.0015954396,0.003169897,0.014385675,0.012561772,0.010030659,0.010771761],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0048609367,0.0026792747,0.013063703,0.0038744474,0.0014900533,0.0010985931,0.0049124113,0.05180548,0.02016328,0.016592953,0.43460116,0.44485775],"study_design_scores_gemma":[0.0025144704,0.0019837676,0.009885616,0.0011673006,0.0004393173,0.0016489167,0.004788653,0.6652326,0.043469623,0.10644496,0.16168447,0.0007403281],"about_ca_topic_score_codex":0.011129306,"about_ca_topic_score_gemma":0.016016804,"teacher_disagreement_score":0.0500135,"about_ca_system_score_codex":0.002951848,"about_ca_system_score_gemma":0.0050049955,"threshold_uncertainty_score":0.26449984},"labels":[],"label_agreement":null},{"id":"W3114178959","doi":"10.18653/v1/2020.aacl-main.51","title":"Systematically Exploring Redundancy Reduction in Summarizing Long Documents","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Redundancy (engineering); Automatic summarization; Computer science; Context (archaeology); Artificial intelligence; Information retrieval; Natural language processing; Data mining","score_opus":0.09334047331948332,"score_gpt":0.26549046444250096,"score_spread":0.17214999112301765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3114178959","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14909124,0.023805765,0.80894315,0.0019921747,0.0006142117,0.00056533975,0.005220441,0.005978072,0.0037896323],"genre_scores_gemma":[0.33444095,0.0042106668,0.6353054,0.00040635988,0.00086564076,0.00046619697,0.020298542,0.0008914245,0.0031148612],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973999,0.0011809816,0.00024391183,0.0005147342,0.00049965916,0.00016088181],"domain_scores_gemma":[0.98959875,0.0068142964,0.000533265,0.0010257751,0.0017939697,0.0002339445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003207931,0.0019561667,0.0023704232,0.0068400097,0.0013310544,0.0021345182,0.0020601784,0.0012659922,0.0020342292],"category_scores_gemma":[0.013146698,0.0007499537,0.0014118725,0.005613565,0.000731533,0.003829468,0.0019146731,0.0011792604,0.001670064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010185854,0.00050136744,0.0055037644,0.0025600179,0.0007551671,0.0006196561,0.0014718875,0.03764346,0.032447223,0.0071526123,0.04092339,0.86940277],"study_design_scores_gemma":[0.00028580803,0.0013303299,0.0065736636,0.00052327756,0.0017258916,0.00085722917,0.0023347528,0.8144544,0.039791904,0.09341091,0.038530644,0.00018116442],"about_ca_topic_score_codex":0.0027260545,"about_ca_topic_score_gemma":0.0051787323,"teacher_disagreement_score":0.0068400097,"about_ca_system_score_codex":0.0005124386,"about_ca_system_score_gemma":0.0016922744,"threshold_uncertainty_score":0.01696539},"labels":[],"label_agreement":null},{"id":"W3114312194","doi":"10.18653/v1/2020.inlg-1.11","title":"Towards Generating Query to Perform Query Focused Abstractive Summarization using Pre-trained Model","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada; University of Lethbridge","keywords":"Computer science; Automatic summarization; Query expansion; Natural language processing; Sentence; Artificial intelligence; Information retrieval; Web search query; Query optimization; RDF query language; Query language; Web query classification; Search engine","score_opus":0.06164828382117375,"score_gpt":0.28109674329067325,"score_spread":0.21944845946949948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3114312194","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04368653,0.0034307532,0.91401345,0.0013848614,0.00031939152,0.00089405605,0.0055311127,0.027387103,0.003352704],"genre_scores_gemma":[0.25300705,0.0017043125,0.696206,0.0012091092,0.0006069405,0.0010508691,0.03391663,0.00082729134,0.011471737],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99912316,0.00023513439,0.00008107001,0.00030360563,0.0001709183,0.000086046035],"domain_scores_gemma":[0.9981894,0.00076086796,0.00011265321,0.00021722945,0.000630887,0.00008889295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014304816,0.0022206805,0.0012290013,0.0019437568,0.0004357317,0.0012383007,0.0015154097,0.0014022512,0.0031089715],"category_scores_gemma":[0.0041300557,0.00036158419,0.0013239105,0.0014073793,0.00040179575,0.0023295172,0.0008308791,0.0017649897,0.0032009522],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000860118,0.00066020247,0.0027904431,0.0013409042,0.00027271584,0.00045650842,0.00072710216,0.08571011,0.088321246,0.0062124976,0.07484853,0.7377996],"study_design_scores_gemma":[0.0001357623,0.00050787226,0.0011565465,0.000042987303,0.0001600188,0.00019353542,0.00028883427,0.9447839,0.027716527,0.005503677,0.019452745,0.000057539834],"about_ca_topic_score_codex":0.009072704,"about_ca_topic_score_gemma":0.012534263,"teacher_disagreement_score":0.009072704,"about_ca_system_score_codex":0.0010613919,"about_ca_system_score_gemma":0.0022839385,"threshold_uncertainty_score":0.018039763},"labels":[],"label_agreement":null},{"id":"W3114513193","doi":"10.1145/3437963.3441728","title":"CalibreNet: Calibration Networks for Multilingual Sequence Labeling","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Natural Science Foundation of China; Jilin Scientific and Technological Development Program; Jilin Province Development and Reform Commission","keywords":"Computer science; Benchmark (surveying); Sequence labeling; Task (project management); Artificial intelligence; Phrase; Sequence (biology); Natural language processing; Named-entity recognition; Obstacle; Boundary (topology); Resource (disambiguation); Machine learning; Mathematics","score_opus":0.05275442905890881,"score_gpt":0.29703419629940675,"score_spread":0.24427976724049794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3114513193","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0145138865,0.0007700353,0.954894,0.00035187442,0.00017036108,0.00020470057,0.001783797,0.022832416,0.0044790115],"genre_scores_gemma":[0.24133526,0.00070832926,0.7152481,0.0012097814,0.00023453799,0.0009959013,0.022292731,0.0027079307,0.015267341],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985763,0.0003637837,0.000056392917,0.0006815499,0.0001936093,0.00012836102],"domain_scores_gemma":[0.9972234,0.0012032049,0.00018021646,0.000656667,0.0006143371,0.00012223875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018985287,0.0022700173,0.0010489214,0.0018977545,0.0011867462,0.0013051553,0.0032926851,0.0028680805,0.007962229],"category_scores_gemma":[0.008947108,0.00090798433,0.00121767,0.0018108401,0.0009811962,0.0045944895,0.0032675532,0.0038944634,0.0056027737],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046538058,0.00031914146,0.00309178,0.00037884206,0.00018403622,0.0003838209,0.000719839,0.13442405,0.017876111,0.015465812,0.050552867,0.77613837],"study_design_scores_gemma":[0.000035510137,0.000094036564,0.00055154593,0.000065300555,0.000037935188,0.00012880159,0.00016261548,0.9482547,0.008702776,0.027454594,0.014471364,0.000040842766],"about_ca_topic_score_codex":0.006883116,"about_ca_topic_score_gemma":0.013115203,"teacher_disagreement_score":0.007962229,"about_ca_system_score_codex":0.0011904587,"about_ca_system_score_gemma":0.0017777982,"threshold_uncertainty_score":0.026636302},"labels":[],"label_agreement":null},{"id":"W3115018051","doi":"10.3390/risks9010007","title":"Mining Actuarial Risk Predictors in Accident Descriptions Using Recurrent Neural Networks","year":2020,"lang":"en","type":"article","venue":"Risks","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Accident (philosophy); Computer science; Artificial neural network; Poisson regression; Regression; Artificial intelligence; Machine learning; Task (project management); Representation (politics); Profit (economics); Data mining; Econometrics; Statistics; Mathematics; Engineering","score_opus":0.12687238660100658,"score_gpt":0.3084885536545194,"score_spread":0.1816161670535128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3115018051","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49476472,0.00081740756,0.49901104,0.0007279664,0.000054093616,0.000076193406,0.0016790598,0.001303272,0.0015662488],"genre_scores_gemma":[0.9600941,0.00023929015,0.0365663,0.000038507926,0.000052402855,0.000051361603,0.0018220962,0.000030432486,0.0011055073],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99947673,0.00016284348,0.000048325073,0.00014119214,0.000118409036,0.000052500865],"domain_scores_gemma":[0.9966633,0.0023108693,0.0005268379,0.00019373144,0.00024422194,0.000060965365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012171848,0.0007661234,0.0005533187,0.0021114915,0.00023991069,0.00065481145,0.0009362119,0.0006939,0.00073340995],"category_scores_gemma":[0.0061801192,0.00046223972,0.0007552764,0.001453143,0.0003036825,0.0012691867,0.00057078636,0.0011382396,0.00032762677],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022285449,0.00022870049,0.04124647,0.00013348341,0.00018077623,0.0005072771,0.00030725924,0.81329155,0.0026975079,0.0062430585,0.002843767,0.13209715],"study_design_scores_gemma":[0.0000020947746,0.000008616698,0.001246327,0.0000050725293,0.000009401307,0.000016083319,0.000017216124,0.99580646,0.00032259693,0.0023817485,0.00017974703,0.0000045301963],"about_ca_topic_score_codex":0.0068756426,"about_ca_topic_score_gemma":0.010491457,"teacher_disagreement_score":0.0068756426,"about_ca_system_score_codex":0.00069633016,"about_ca_system_score_gemma":0.00048598365,"threshold_uncertainty_score":0.013671219},"labels":[],"label_agreement":null},{"id":"W3115355887","doi":"10.18653/v1/2020.aacl-main.63","title":"Improving Context Modeling in Neural Topic Segmentation","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Western University","funders":"","keywords":"Computer science; Robustness (evolution); Artificial intelligence; Machine learning; Segmentation; Artificial neural network; Key (lock); Context (archaeology); Context model; Task (project management); Transfer of learning; German; Natural language processing","score_opus":0.06253534279808755,"score_gpt":0.253566698615409,"score_spread":0.19103135581732145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3115355887","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07587623,0.014248385,0.8922482,0.0011463288,0.0006551107,0.00015824413,0.0008864882,0.0102984095,0.004482645],"genre_scores_gemma":[0.66840416,0.003487984,0.31273836,0.0006164527,0.0010620168,0.00034247607,0.0046152896,0.0018501548,0.006883051],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99886036,0.00042236721,0.00006846379,0.0003785158,0.0001145308,0.00015575282],"domain_scores_gemma":[0.9975388,0.0015938866,0.000086484186,0.0002945052,0.00034779753,0.00013851261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020526883,0.0015725708,0.0021329953,0.0021555892,0.0012763288,0.0021599722,0.0019922708,0.0021251603,0.003130754],"category_scores_gemma":[0.0071212566,0.0008904698,0.0016155211,0.002254582,0.00056915823,0.0040673474,0.0020599195,0.0026850302,0.0026052175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010055581,0.00031252354,0.004713376,0.00044845033,0.00038054225,0.00017312962,0.0006580563,0.10534752,0.012378493,0.011285874,0.03235458,0.8309419],"study_design_scores_gemma":[0.00005444062,0.000055322253,0.0008690246,0.00004086773,0.00012029165,0.000059752107,0.00010222445,0.9763084,0.0038134987,0.014917749,0.0036349588,0.000023362047],"about_ca_topic_score_codex":0.013806508,"about_ca_topic_score_gemma":0.024925118,"teacher_disagreement_score":0.013806508,"about_ca_system_score_codex":0.0009627793,"about_ca_system_score_gemma":0.0014719478,"threshold_uncertainty_score":0.02745229},"labels":[],"label_agreement":null},{"id":"W3115763169","doi":"10.18653/v1/2020.coling-main.456","title":"Explain by Evidence: An Explainable Memory-based Neural Network for Question Answering","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Interpretability; Computer science; Artificial intelligence; TRACE (psycholinguistics); Machine learning; Tracing; Artificial neural network; Question answering","score_opus":0.059165063389809085,"score_gpt":0.28928172777272876,"score_spread":0.23011666438291967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3115763169","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10165097,0.003312293,0.8822848,0.0027706246,0.00019624546,0.00015237171,0.0015453221,0.0040964363,0.003990896],"genre_scores_gemma":[0.83999157,0.00096398714,0.15153748,0.0006844531,0.00012716131,0.00020221304,0.0016627859,0.00012519823,0.004705171],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997948,0.000048373287,0.00001282212,0.000082426166,0.000034082932,0.000027506127],"domain_scores_gemma":[0.9991055,0.0005044988,0.00010456224,0.00015387165,0.00009664286,0.000034935463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007385144,0.0007933162,0.00057181146,0.00071168126,0.0003188067,0.000901692,0.0020537435,0.001425123,0.0024047992],"category_scores_gemma":[0.003720003,0.00039288134,0.00078252814,0.0006929956,0.0006211447,0.0022890072,0.0012591142,0.0019721873,0.00047307464],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055770465,0.00025910497,0.0065575168,0.00031285876,0.00033037222,0.00029523263,0.0005318228,0.37728995,0.009378595,0.02855311,0.0092216795,0.56671214],"study_design_scores_gemma":[0.000014881267,0.000045684046,0.0005140512,0.000025927835,0.000055581542,0.000035721394,0.000019300849,0.974766,0.0015066832,0.021744745,0.0012584163,0.000012973583],"about_ca_topic_score_codex":0.0073301927,"about_ca_topic_score_gemma":0.013726747,"teacher_disagreement_score":0.0073301927,"about_ca_system_score_codex":0.0008505535,"about_ca_system_score_gemma":0.0007548123,"threshold_uncertainty_score":0.014575005},"labels":[],"label_agreement":null},{"id":"W3115808866","doi":"10.18653/v1/2020.coling-main.58","title":"TIMBERT: Toponym Identifier For The Medical Domain Based on BERT","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Sentence; Identifier; Domain (mathematical analysis); Natural language processing; Set (abstract data type); Artificial intelligence; Task (project management); Test set; Identification (biology); Process (computing); Named-entity recognition; Programming language","score_opus":0.031082334903676383,"score_gpt":0.26768369044451207,"score_spread":0.2366013555408357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3115808866","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049815472,0.001405821,0.8725043,0.0013106116,0.0007774898,0.000495,0.018428769,0.049172193,0.0060903723],"genre_scores_gemma":[0.3935353,0.0008348588,0.5617974,0.0005349647,0.00039857437,0.00047690963,0.034178346,0.0016308806,0.0066127656],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989492,0.00030598667,0.000089630645,0.00033429934,0.00025763418,0.000063244326],"domain_scores_gemma":[0.99662954,0.00190596,0.00034685706,0.00046383322,0.00049170665,0.00016216293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001742403,0.0013280704,0.00048120294,0.0023465704,0.00060930144,0.0010663723,0.0011530722,0.0012009982,0.005410635],"category_scores_gemma":[0.005075821,0.00047415952,0.0011061308,0.0011765592,0.00037862363,0.004352855,0.0014254269,0.0017640035,0.0034021533],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013824188,0.0003432345,0.016606567,0.0021699204,0.00052453834,0.0015168154,0.0018169092,0.06337288,0.05673615,0.034257077,0.20920749,0.6120659],"study_design_scores_gemma":[0.00008221811,0.0003891589,0.007363233,0.00015019719,0.00015780705,0.0020866601,0.0004630442,0.8319927,0.027416259,0.03356926,0.09615946,0.00017004627],"about_ca_topic_score_codex":0.0035856888,"about_ca_topic_score_gemma":0.007076901,"teacher_disagreement_score":0.005410635,"about_ca_system_score_codex":0.000748759,"about_ca_system_score_gemma":0.0011494834,"threshold_uncertainty_score":0.01810038},"labels":[],"label_agreement":null},{"id":"W3116083993","doi":"10.1609/aaai.v35i15.17627","title":"Learning Contextual Representations for Semantic Parsing with Generation-Augmented Pre-Training","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language processing; Utterance; Artificial intelligence; Parsing; Leverage (statistics); Schema (genetic algorithms); SQL; Language model; Information retrieval; Programming language","score_opus":0.15111855803125188,"score_gpt":0.331868255043637,"score_spread":0.1807496970123851,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3116083993","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.074059814,0.0009637151,0.8864401,0.00073125924,0.00015013547,0.00021100635,0.0022953886,0.031971864,0.0031765925],"genre_scores_gemma":[0.5316002,0.0004263009,0.44721952,0.00076466973,0.00009455292,0.0005598839,0.0150497835,0.0013946125,0.0028904886],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896896,0.00042629227,0.000046310124,0.0003600055,0.00010208521,0.00009639945],"domain_scores_gemma":[0.9971615,0.0017874589,0.00009862869,0.00060967775,0.0002732773,0.000069441274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015223157,0.0016910913,0.0007848316,0.0008336882,0.0005226561,0.0008366901,0.002103564,0.0012765737,0.004466098],"category_scores_gemma":[0.005753723,0.00071705296,0.00107611,0.0010018508,0.000791762,0.003025438,0.0016860475,0.0031869537,0.0018593541],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052002235,0.0005345633,0.005356467,0.00048807292,0.0001877165,0.00045512017,0.0007177239,0.37865898,0.016996084,0.010956855,0.031118805,0.5540097],"study_design_scores_gemma":[0.000031155912,0.000081946535,0.000540661,0.00002339437,0.000038001144,0.00006157292,0.00009011119,0.97646654,0.0076659806,0.011240388,0.0037430336,0.000017181383],"about_ca_topic_score_codex":0.00649394,"about_ca_topic_score_gemma":0.01714286,"teacher_disagreement_score":0.00649394,"about_ca_system_score_codex":0.0010177798,"about_ca_system_score_gemma":0.0017494173,"threshold_uncertainty_score":0.01494056},"labels":[],"label_agreement":null},{"id":"W3116453230","doi":"10.18653/v1/2020.coling-main.106","title":"Learning Efficient Task-Specific Meta-Embeddings with Word Prisms","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute; Bank of Canada","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Word (group theory); Embedding; Inference; Word embedding; Context (archaeology); Task (project management); Set (abstract data type); Artificial intelligence; Natural language processing; Mathematics","score_opus":0.04411524210167301,"score_gpt":0.22848442374406747,"score_spread":0.18436918164239446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3116453230","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03852422,0.0010644884,0.9451109,0.00034865356,0.00018806673,0.00012123189,0.0013081692,0.011844236,0.001490003],"genre_scores_gemma":[0.467493,0.0009147181,0.5145769,0.00032067578,0.0001732247,0.00046810077,0.009588397,0.0013037965,0.0051611806],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987901,0.0003870698,0.00012867853,0.00039778897,0.000198168,0.00009822385],"domain_scores_gemma":[0.99750465,0.0009337275,0.00019869288,0.00088942045,0.00036534655,0.000108267755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001961432,0.0023428013,0.001152577,0.0013007964,0.00031244717,0.0015684792,0.00223378,0.0014549362,0.0028153157],"category_scores_gemma":[0.0071282256,0.00084921933,0.0018795184,0.0019466201,0.0006841322,0.0077868337,0.003272648,0.0030084548,0.0023973717],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064184197,0.00043131525,0.0047228467,0.0008301645,0.0004523955,0.00025821535,0.00045303858,0.20288603,0.020739775,0.016132116,0.018265793,0.7341866],"study_design_scores_gemma":[0.00007493018,0.00017794926,0.0005435469,0.000043002612,0.000081392434,0.00013614832,0.000114463524,0.95694166,0.009587589,0.028098388,0.004166152,0.00003464246],"about_ca_topic_score_codex":0.0015467817,"about_ca_topic_score_gemma":0.003379746,"teacher_disagreement_score":0.0028153157,"about_ca_system_score_codex":0.0005873003,"about_ca_system_score_gemma":0.0013243594,"threshold_uncertainty_score":0.010373116},"labels":[],"label_agreement":null},{"id":"W3116598976","doi":"10.5281/zenodo.4513822","title":"A Methodology for Hierarchical Classification of Semantic Answer Types of Questions","year":2020,"lang":"en","type":"article","venue":"Research Publications (Maastricht University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Information retrieval","score_opus":0.3338115113383785,"score_gpt":0.3907639597184681,"score_spread":0.056952448380089615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3116598976","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008996537,0.00046625896,0.9760477,0.0004722007,0.00017988558,0.0009646614,0.0040988396,0.0062019057,0.0025720356],"genre_scores_gemma":[0.08874403,0.000188674,0.89411736,0.00026538115,0.00017625774,0.0014487987,0.010195764,0.00030696526,0.004556657],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9953259,0.0008661174,0.0005828912,0.0013622594,0.001465965,0.00039685884],"domain_scores_gemma":[0.9924005,0.0032551133,0.00062062655,0.0010740188,0.0024015184,0.00024810896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045898315,0.0014926561,0.0011518745,0.010484322,0.0017396947,0.0022625518,0.0030769096,0.0023772223,0.0055703367],"category_scores_gemma":[0.015068854,0.0005845429,0.0028700326,0.006096204,0.00095382385,0.0032055993,0.0016962141,0.0031949044,0.0039035399],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041436983,0.00067016325,0.012714439,0.00055385125,0.00022646469,0.00020227693,0.0012556817,0.013275543,0.017252387,0.032328952,0.036669478,0.8844364],"study_design_scores_gemma":[0.00013279362,0.00031828112,0.0130314175,0.00021341865,0.00019122192,0.0006354138,0.0008366746,0.7797198,0.024218062,0.12301854,0.057557475,0.0001268285],"about_ca_topic_score_codex":0.0126023935,"about_ca_topic_score_gemma":0.018038236,"teacher_disagreement_score":0.0126023935,"about_ca_system_score_codex":0.0024103092,"about_ca_system_score_gemma":0.0028751653,"threshold_uncertainty_score":0.02505809},"labels":[],"label_agreement":null},{"id":"W3116946603","doi":"10.2196/22795","title":"Adapting Bidirectional Encoder Representations from Transformers (BERT) to Assess Clinical Semantic Textual Similarity: Algorithm Development and Validation Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Deutschen Konsortium für Translationale Krebsforschung; Deutsches Krebsforschungszentrum","keywords":"Computer science; Encoder; Pearson product-moment correlation coefficient; Leverage (statistics); Artificial intelligence; Semantic similarity; Natural language processing; Relationship extraction; Sentence; Test set; Ground truth; Machine learning; Information extraction; Statistics","score_opus":0.1256573497326064,"score_gpt":0.38201198221689464,"score_spread":0.25635463248428825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3116946603","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28855497,0.0021168296,0.68284684,0.000656785,0.0003315168,0.00082451093,0.0020456556,0.018009918,0.004612993],"genre_scores_gemma":[0.68426174,0.0004813521,0.3045365,0.00026629653,0.000086371496,0.0005015462,0.006316203,0.0004383204,0.0031117443],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985139,0.0006410738,0.000104703424,0.0003041671,0.0002935447,0.0001426194],"domain_scores_gemma":[0.99421436,0.0039861365,0.00021514772,0.0003974481,0.0010163548,0.00017062132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038126425,0.0015971283,0.0009923778,0.0017578756,0.00040849322,0.0010768289,0.0015429336,0.0014258702,0.00278423],"category_scores_gemma":[0.013821286,0.00038058034,0.00075597956,0.0010131163,0.000410931,0.0013259491,0.0014161505,0.0016907856,0.001312362],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007723265,0.0005432544,0.010066799,0.00028483965,0.00024297465,0.00016764474,0.00014797955,0.2762241,0.0041339244,0.0028038742,0.010985026,0.6936272],"study_design_scores_gemma":[0.000046049267,0.00012450757,0.0007761682,0.000018388122,0.000024270887,0.00006141814,0.00004605959,0.994561,0.0025135288,0.0012015518,0.0006162555,0.000010821672],"about_ca_topic_score_codex":0.013812566,"about_ca_topic_score_gemma":0.014051289,"teacher_disagreement_score":0.013812566,"about_ca_system_score_codex":0.0013581864,"about_ca_system_score_gemma":0.0021922658,"threshold_uncertainty_score":0.02746433},"labels":[],"label_agreement":null},{"id":"W3117829624","doi":"10.48550/arxiv.2012.09936","title":"Named Entity Recognition in the Legal Domain using a Pointer Generator Network","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Named-entity recognition; Computer science; Pointer (user interface); Security token; Natural language processing; Task (project management); Artificial intelligence; Annotation; Information retrieval; Computer security","score_opus":0.12920337610299756,"score_gpt":0.19757034710313318,"score_spread":0.06836697100013561,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3117829624","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10915703,0.0006172194,0.87336254,0.0013579591,0.00019274656,0.00018112615,0.0010053285,0.007559154,0.0065668994],"genre_scores_gemma":[0.61584294,0.00050412567,0.3645454,0.0005170807,0.00015425442,0.00028951903,0.004216098,0.00028246973,0.013648108],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99969995,0.00008009734,0.000015263173,0.00012507645,0.00004822058,0.000031405332],"domain_scores_gemma":[0.99894553,0.0006250558,0.00006127232,0.00014761089,0.00018421016,0.000036268942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010058173,0.00052865874,0.0004097551,0.0008643241,0.0005097344,0.00073331845,0.0012811915,0.0012909865,0.0029388992],"category_scores_gemma":[0.0033750066,0.00031359613,0.0006247998,0.0009151585,0.0005723581,0.0023859218,0.0009817535,0.0013456271,0.0013889022],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023573193,0.0003241229,0.003001838,0.00014310326,0.00008858873,0.0005946968,0.00025024824,0.4976459,0.015714511,0.018662812,0.014885687,0.44845274],"study_design_scores_gemma":[0.0000073860797,0.000024077523,0.00027028177,0.0000073084025,0.00000883018,0.000037289366,0.0000150511,0.9895359,0.003128806,0.0059011346,0.0010570303,0.000006854794],"about_ca_topic_score_codex":0.0055951835,"about_ca_topic_score_gemma":0.007967469,"teacher_disagreement_score":0.0055951835,"about_ca_system_score_codex":0.0008389186,"about_ca_system_score_gemma":0.0007095994,"threshold_uncertainty_score":0.011125207},"labels":[],"label_agreement":null},{"id":"W3118205297","doi":"10.18653/v1/2020.coling-main.455","title":"Intra-Correlation Encoding for Chinese Sentence Intention Matching","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Regional Municipality of Niagara; Brock University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Sentence; Natural language processing; Artificial intelligence; Encoding (memory); Ambiguity; Word (group theory); Matching (statistics); Embedding; Granularity; Feature (linguistics); Linguistics; Mathematics","score_opus":0.029475001894958226,"score_gpt":0.2619573841806866,"score_spread":0.23248238228572837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118205297","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.098236434,0.0019912352,0.8820794,0.00051158614,0.00023180636,0.0003386246,0.0025765623,0.0072214017,0.00681295],"genre_scores_gemma":[0.7630703,0.00086862996,0.22095181,0.00029554107,0.00020122922,0.00045381993,0.006405489,0.00041687177,0.00733631],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994318,0.00014019373,0.000060398615,0.0001622386,0.0001321079,0.00007327285],"domain_scores_gemma":[0.9992009,0.00023451263,0.00009330128,0.00014606187,0.00028103965,0.000044309076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00077210803,0.0007878011,0.0006094397,0.0020496722,0.00033773898,0.00057169615,0.0008315955,0.0004563235,0.0031021305],"category_scores_gemma":[0.0027718025,0.00020042065,0.00083313737,0.0019682092,0.00029214,0.0017720661,0.0007883898,0.000827624,0.001441741],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005671337,0.00035122488,0.006440041,0.00034076153,0.00011162354,0.0002624614,0.00056915823,0.03341783,0.021556605,0.014238635,0.023603918,0.8985405],"study_design_scores_gemma":[0.000036709742,0.00016066195,0.004448393,0.000036090772,0.0001011686,0.00016418842,0.00013670344,0.96005476,0.012988894,0.013019018,0.008807545,0.000045865087],"about_ca_topic_score_codex":0.009267111,"about_ca_topic_score_gemma":0.010708367,"teacher_disagreement_score":0.009267111,"about_ca_system_score_codex":0.0007034014,"about_ca_system_score_gemma":0.0013497293,"threshold_uncertainty_score":0.018426359},"labels":[],"label_agreement":null},{"id":"W3118212635","doi":"10.18653/v1/2020.coling-main.147","title":"Unsupervised Fact Checking by Counter-Weighted Positive and Negative Evidential Paths in A Knowledge Graph","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea","keywords":"Computer science; Statement (logic); Misinformation; Knowledge graph; Artificial intelligence; Information retrieval; Computer security; Political science; Law","score_opus":0.021901528073056585,"score_gpt":0.24976124317122328,"score_spread":0.2278597150981667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118212635","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22681275,0.0016954711,0.747234,0.002124466,0.00020513538,0.0005665191,0.008148097,0.0062321085,0.0069814986],"genre_scores_gemma":[0.6344167,0.0005353454,0.34614074,0.0004919186,0.00011821192,0.0001759771,0.014594655,0.00034574952,0.0031806452],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99611354,0.00068810064,0.00034839797,0.0015353813,0.0010890178,0.00022548002],"domain_scores_gemma":[0.97803086,0.013939436,0.002190737,0.0022862875,0.0032031653,0.000349477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023617765,0.0010092744,0.0006275023,0.006496626,0.0015036415,0.0021805419,0.0021061292,0.0017956054,0.0016614177],"category_scores_gemma":[0.02149735,0.00052339915,0.0014048083,0.0027394367,0.0016807409,0.00496809,0.0016252177,0.0021178015,0.0006371985],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009476648,0.0004933121,0.081493735,0.0014075392,0.0006885944,0.0045651793,0.0026912065,0.18828729,0.02050251,0.06347534,0.03103546,0.60441214],"study_design_scores_gemma":[0.00007234213,0.00009898593,0.009764606,0.0002610888,0.00042564713,0.001029907,0.0006639307,0.7943299,0.019224634,0.15081228,0.0232372,0.00007947109],"about_ca_topic_score_codex":0.018365592,"about_ca_topic_score_gemma":0.0399312,"teacher_disagreement_score":0.018365592,"about_ca_system_score_codex":0.0016689848,"about_ca_system_score_gemma":0.0027680974,"threshold_uncertainty_score":0.03651738},"labels":[],"label_agreement":null},{"id":"W3118454345","doi":"","title":"On-the-Fly Attention Modularization for Neural Generation","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Fluency; Computer science; Modular programming; Artificial intelligence; Artificial neural network; Sentence; Natural language processing; Inference; Cognitive psychology; Machine learning; Psychology; Programming language","score_opus":0.12865905516853385,"score_gpt":0.18980302845342964,"score_spread":0.06114397328489579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118454345","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03296469,0.00019375532,0.9610924,0.00025505794,0.000039257007,0.0000531351,0.000106199535,0.0034601383,0.0018353064],"genre_scores_gemma":[0.68836236,0.00019419064,0.3074568,0.00021849376,0.00008330741,0.00016698966,0.00035835482,0.0005430414,0.0026164323],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99963725,0.00011834935,0.00002324588,0.000105889965,0.00007053234,0.000044691275],"domain_scores_gemma":[0.9989649,0.00048774437,0.000069717826,0.00031859466,0.000100937905,0.000058147965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00077905675,0.00058733724,0.0004872597,0.00044176742,0.00027651805,0.0006041462,0.0013033076,0.000670983,0.0035104959],"category_scores_gemma":[0.003663364,0.0003337694,0.0007329363,0.0003996988,0.0006919532,0.001750061,0.001615745,0.001335537,0.0008948364],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025233853,0.00018929262,0.0018520284,0.00023809797,0.00014869498,0.00022189744,0.00041486372,0.3151971,0.04889339,0.05497792,0.0063632703,0.5712511],"study_design_scores_gemma":[0.00001971066,0.00003012904,0.00018129442,0.00000727134,0.000019865405,0.00002998666,0.000015438325,0.952054,0.007887152,0.03848817,0.0012600654,0.000006931846],"about_ca_topic_score_codex":0.0012333556,"about_ca_topic_score_gemma":0.0025543303,"teacher_disagreement_score":0.0035104959,"about_ca_system_score_codex":0.0006676879,"about_ca_system_score_gemma":0.0007139865,"threshold_uncertainty_score":0.011743784},"labels":[],"label_agreement":null},{"id":"W3119437245","doi":"10.18653/v1/2021.acl-long.519","title":"End-to-End Training of Neural Retrievers for Open-Domain Question Answering","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Canadian Institute for Advanced Research; Nvidia","keywords":"Computer science; Artificial intelligence; Question answering; Context (archaeology); Machine learning; Domain (mathematical analysis); Task (project management); Artificial neural network; Biomedical text mining; Natural language processing; Text mining; Mathematics","score_opus":0.07985047469528585,"score_gpt":0.3197809994370351,"score_spread":0.23993052474174925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119437245","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4181784,0.013843721,0.46920824,0.0018380521,0.001571359,0.0008759156,0.004950011,0.07375867,0.015775664],"genre_scores_gemma":[0.688682,0.001429763,0.26572338,0.0007028126,0.0004879503,0.00063531386,0.021010553,0.0015180615,0.019810088],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99892175,0.0002661848,0.00008693443,0.00040250013,0.00014724773,0.00017536066],"domain_scores_gemma":[0.99677926,0.0017445004,0.000108189124,0.00046645728,0.000713018,0.00018859559],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003050351,0.0022368685,0.0017431829,0.0017114296,0.0009939849,0.0017819632,0.0040334603,0.003499349,0.0077378266],"category_scores_gemma":[0.008415372,0.001106127,0.0013216994,0.001365301,0.00057721353,0.0031676409,0.0021607452,0.0032185365,0.008755799],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013067046,0.00096916757,0.0037549455,0.00035129773,0.0003950837,0.00025625815,0.00038956528,0.05353952,0.014466407,0.0014011109,0.040701654,0.8824684],"study_design_scores_gemma":[0.00017575863,0.00032595298,0.0017895651,0.000042690863,0.00012428816,0.00011685594,0.00023360882,0.97638583,0.014242602,0.0027710507,0.0037584014,0.000033465836],"about_ca_topic_score_codex":0.010563907,"about_ca_topic_score_gemma":0.019048912,"teacher_disagreement_score":0.010563907,"about_ca_system_score_codex":0.0012754593,"about_ca_system_score_gemma":0.0011924476,"threshold_uncertainty_score":0.025885582},"labels":[],"label_agreement":null},{"id":"W3119519251","doi":"10.1609/aaai.v35i16.17681","title":"What's the Best Place for an AI Conference, Vancouver or _______: Why Completing Comparative Questions is Difficult","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Competence (human resources); Artificial intelligence; sort; Natural language processing; Benchmark (surveying); Set (abstract data type); Task (project management); Language model; Machine learning; Cognitive science; Psychology; Information retrieval","score_opus":0.20765916333050938,"score_gpt":0.36477807941224205,"score_spread":0.15711891608173267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119519251","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1865802,0.011196064,0.027711945,0.09881852,0.0062001743,0.00042472532,0.022103913,0.007116349,0.6398482],"genre_scores_gemma":[0.6952793,0.005767223,0.043406542,0.0051122108,0.0008747542,0.0002214436,0.02562975,0.001252831,0.22245595],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993579,0.0002195016,0.000030905325,0.00015697203,0.0001302013,0.00010456522],"domain_scores_gemma":[0.99753004,0.00066305394,0.00012824059,0.00020783328,0.00078048796,0.00069047505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013002687,0.00056387635,0.00041393752,0.0007662399,0.0032267051,0.0045204894,0.0007929296,0.0015209964,0.04888696],"category_scores_gemma":[0.007582098,0.00025518617,0.00033317774,0.0014677515,0.00096611807,0.0035302844,0.0009270457,0.0021529284,0.01137626],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047451715,0.00015789612,0.014469725,0.00065984274,0.000083050414,0.0006339146,0.0026890626,0.004381315,0.0033709554,0.02303216,0.6434322,0.30661538],"study_design_scores_gemma":[0.00006335695,0.000088167086,0.028612195,0.00049470627,0.000054176468,0.0004165666,0.0076354565,0.012076243,0.004215505,0.029315904,0.9169004,0.00012744415],"about_ca_topic_score_codex":0.17785506,"about_ca_topic_score_gemma":0.41248897,"teacher_disagreement_score":0.17785506,"about_ca_system_score_codex":0.0037287928,"about_ca_system_score_gemma":0.0038791772,"threshold_uncertainty_score":0.35363966},"labels":[],"label_agreement":null},{"id":"W3119636502","doi":"10.18653/v1/2021.acl-long.569","title":"UnNatural Language Inference","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University","funders":"McGill University","keywords":"Computer science; Inference; Natural language processing; Artificial intelligence; Mandarin Chinese; Syntax; Suite; Transformer; Language model; Word order; Linguistics; History","score_opus":0.029339431725579488,"score_gpt":0.29680368723324707,"score_spread":0.26746425550766756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119636502","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005240675,0.0038473054,0.93364716,0.0070865396,0.0014747287,0.000065100045,0.0011905486,0.0034209033,0.044026952],"genre_scores_gemma":[0.344752,0.00492878,0.5878454,0.0042970576,0.00314662,0.00033710865,0.006944004,0.0037335085,0.044015456],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963182,0.001819338,0.00017412174,0.0008417296,0.0007164472,0.00013020387],"domain_scores_gemma":[0.9910657,0.0054239994,0.00016592571,0.002505131,0.0006893514,0.00014990939],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041651423,0.0010618066,0.0009860144,0.0020780882,0.0018504437,0.0036689462,0.00200812,0.0014821112,0.016007125],"category_scores_gemma":[0.01572675,0.0008523124,0.0019163506,0.0016198864,0.0028835917,0.009123254,0.0046518655,0.004371066,0.0053497883],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007742675,0.000060046485,0.0005406354,0.0002889053,0.0001036276,0.00031044183,0.00058401766,0.0036524036,0.0012117584,0.82543784,0.060272485,0.10746041],"study_design_scores_gemma":[0.00001356112,0.0000072098505,0.00011579697,0.00006134121,0.000025784968,0.00018983231,0.00006844634,0.021867178,0.0012289186,0.9148623,0.06154409,0.000015575712],"about_ca_topic_score_codex":0.001720434,"about_ca_topic_score_gemma":0.0030699945,"teacher_disagreement_score":0.016007125,"about_ca_system_score_codex":0.0011320162,"about_ca_system_score_gemma":0.0010809582,"threshold_uncertainty_score":0.05354917},"labels":[],"label_agreement":null},{"id":"W3119707607","doi":"","title":"Joint Training for Learning Cross-lingual Embeddings with Sub-word Information without Parallel Corpora","year":2020,"lang":"en","type":"article","venue":"Joint Conference on Lexical and Computational Semantics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Natural language processing; Word (group theory); Lexicon; Artificial intelligence; Benchmark (surveying); Similarity (geometry); Resource (disambiguation); Joint (building); Linguistics","score_opus":0.09927202779614587,"score_gpt":0.2928176832555031,"score_spread":0.19354565545935726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119707607","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040487327,0.00082587014,0.94999135,0.00024010136,0.00021552762,0.00013453396,0.0005178139,0.005017128,0.0025704016],"genre_scores_gemma":[0.39440602,0.0006998329,0.58454865,0.000512101,0.00023278715,0.00074957055,0.00958028,0.001329477,0.007941328],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99781835,0.00074658083,0.00018953942,0.0008346808,0.0002591512,0.0001516635],"domain_scores_gemma":[0.9958482,0.0017872596,0.0001812878,0.0011972488,0.000830165,0.00015587348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030517161,0.0021797046,0.0013993684,0.0019090053,0.0007704877,0.0016822569,0.0020766677,0.001520788,0.005011544],"category_scores_gemma":[0.010063873,0.0009860059,0.0014282458,0.0022994701,0.0008997087,0.0068673776,0.0040239473,0.0029804774,0.004888896],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042363038,0.0004981661,0.004489337,0.00036161253,0.00032364667,0.00023586517,0.00041814763,0.07250282,0.014186546,0.00913809,0.012448215,0.88497394],"study_design_scores_gemma":[0.00005658681,0.00022774252,0.0010251019,0.000046903388,0.00008817858,0.00017088665,0.00027333247,0.9537618,0.010024036,0.02534433,0.008940029,0.000040943665],"about_ca_topic_score_codex":0.0023463068,"about_ca_topic_score_gemma":0.004596086,"teacher_disagreement_score":0.005011544,"about_ca_system_score_codex":0.0006053852,"about_ca_system_score_gemma":0.0014148941,"threshold_uncertainty_score":0.016765356},"labels":[],"label_agreement":null},{"id":"W3120832022","doi":"10.18653/v1/2021.emnlp-main.526","title":"Towards Zero-Shot Knowledge Distillation for Natural Language Processing","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Distillation; Task (project management); Benchmark (surveying); Artificial intelligence; Knowledge transfer; Natural language processing; Transfer of learning; Domain knowledge; Variety (cybernetics); Shot (pellet); Machine learning; Domain (mathematical analysis); Knowledge management; Engineering","score_opus":0.08645918045858457,"score_gpt":0.43361835862493486,"score_spread":0.3471591781663503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3120832022","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048416253,0.0013433833,0.93659717,0.0013803798,0.00014632053,0.00012622475,0.00059432496,0.006360498,0.0050355666],"genre_scores_gemma":[0.563806,0.0007706717,0.42135248,0.0010590482,0.00021387181,0.0003093253,0.0029601965,0.0008664651,0.008661884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982198,0.0006653931,0.00007379331,0.000416614,0.0004626081,0.00016171952],"domain_scores_gemma":[0.9959189,0.0026801133,0.00014464023,0.00079106604,0.0003164377,0.00014897465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028041063,0.0014732534,0.0012835445,0.0010310016,0.00085848494,0.001870214,0.0025915857,0.0023007614,0.003460127],"category_scores_gemma":[0.010266993,0.0005637097,0.00096981105,0.0010608813,0.0022122955,0.0058751446,0.0052913562,0.0050153974,0.0018585117],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008926634,0.0005786886,0.0014323447,0.00072885485,0.00016474014,0.00031457175,0.0007015468,0.42160124,0.014236334,0.09028196,0.020902859,0.44816414],"study_design_scores_gemma":[0.00003101061,0.000074786665,0.000090237045,0.000024616886,0.000010643291,0.000048743525,0.00005091702,0.9316118,0.0053439164,0.060348384,0.0023526049,0.0000123802965],"about_ca_topic_score_codex":0.0031935622,"about_ca_topic_score_gemma":0.0052288417,"teacher_disagreement_score":0.003460127,"about_ca_system_score_codex":0.0011542354,"about_ca_system_score_gemma":0.0018328294,"threshold_uncertainty_score":0.014829695},"labels":[],"label_agreement":null},{"id":"W3120950293","doi":"10.18653/v1/2020.aacl-srw.14","title":"Training with Adversaries to Improve Faithfulness of Attention in Neural Machine Translation","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Machine translation; Computer science; Regularization (linguistics); Measure (data warehouse); Translation (biology); Divergence (linguistics); Artificial intelligence; Differentiable function; Machine learning; Quality (philosophy); Trustworthiness; Language model; Natural language processing; Mathematics; Data mining","score_opus":0.04479434594496593,"score_gpt":0.24562341074205213,"score_spread":0.2008290647970862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3120950293","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09568157,0.0025814145,0.8871098,0.0016807754,0.00050383725,0.0001087462,0.00023103786,0.004195848,0.0079069445],"genre_scores_gemma":[0.9162213,0.00047959192,0.0751099,0.0005981205,0.00020933685,0.00013827549,0.00043621974,0.00053919345,0.006268004],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988181,0.0006223579,0.000063017,0.00023962362,0.00013035777,0.00012653605],"domain_scores_gemma":[0.9936458,0.0046507255,0.00019817676,0.001006717,0.00034918863,0.00014940277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002815044,0.001277559,0.0013198289,0.00048304172,0.00067851465,0.00095818593,0.00200656,0.0017472934,0.004077398],"category_scores_gemma":[0.0122173885,0.00071098335,0.00065516075,0.0005998611,0.0016430654,0.0033614344,0.0039333967,0.003405749,0.0016827423],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011070196,0.00032761236,0.0017256031,0.0002692415,0.00017431706,0.00030555486,0.0005059388,0.66016954,0.015867734,0.047420412,0.012964483,0.25916255],"study_design_scores_gemma":[0.000028714068,0.00007178295,0.00010920725,0.000017189357,0.000014994167,0.00003162077,0.00001950367,0.9757494,0.0026772437,0.020647388,0.0006252896,0.000007876619],"about_ca_topic_score_codex":0.00195339,"about_ca_topic_score_gemma":0.0031225572,"teacher_disagreement_score":0.004077398,"about_ca_system_score_codex":0.0007134118,"about_ca_system_score_gemma":0.0008351687,"threshold_uncertainty_score":0.014887512},"labels":[],"label_agreement":null},{"id":"W3122003999","doi":"10.1002/asi.24606","title":"Domain‐topic models with chained dimensions: Charting an emergent domain of a major oncology conference","year":2021,"lang":"en","type":"preprint","venue":"Journal of the Association for Information Science and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Canadian Institutes of Health Research; Agence Nationale de la Recherche; American Society of Clinical Oncology","keywords":"Metadata; Computer science; Domain (mathematical analysis); Context (archaeology); Topic model; Representation (politics); Graph; Cluster analysis; Information retrieval; Dimension (graph theory); Data science; World Wide Web; Artificial intelligence; Theoretical computer science; Mathematics; Geography","score_opus":0.024761635265290613,"score_gpt":0.2767709140196001,"score_spread":0.25200927875430945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122003999","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19544704,0.0029936903,0.7466519,0.0063689705,0.00014501413,0.0002633315,0.0015768787,0.0005999804,0.04595321],"genre_scores_gemma":[0.8290186,0.0013509293,0.16145253,0.00018969436,0.00014686122,0.00036834987,0.0011553467,0.00021295696,0.0061046956],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99843365,0.00083808467,0.00006883467,0.00032639236,0.0002134534,0.00011962124],"domain_scores_gemma":[0.98921657,0.008197346,0.00076390035,0.00072818017,0.0005709902,0.00052296964],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0033011404,0.00036655462,0.00048679358,0.0030128933,0.0014956752,0.0053285495,0.001013837,0.0010672066,0.0053872727],"category_scores_gemma":[0.01598113,0.0004486914,0.00096510607,0.0045599216,0.0028985802,0.008549729,0.0027527919,0.0018744085,0.0008020974],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012857234,0.00005196913,0.012383599,0.00022954535,0.000051931722,0.00019421491,0.009710645,0.057686117,0.0009075747,0.8812135,0.004144319,0.033297945],"study_design_scores_gemma":[0.00004129889,0.00003805829,0.006005914,0.00012867282,0.000038261176,0.00015221692,0.004128934,0.33881733,0.00055340043,0.6112202,0.038816202,0.00005958922],"about_ca_topic_score_codex":0.013781178,"about_ca_topic_score_gemma":0.01244528,"teacher_disagreement_score":0.9969871,"about_ca_system_score_codex":0.003186254,"about_ca_system_score_gemma":0.0019060957,"threshold_uncertainty_score":0.027401924},"labels":[],"label_agreement":null},{"id":"W3122771415","doi":"","title":"Factorizing Declarative and Procedural Knowledge in Structured, Dynamical Environments","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Université de Montréal","funders":"","keywords":"Computer science; Object (grammar); Procedural knowledge; Generalization; Descriptive knowledge; Artificial intelligence; Programming language; Human–computer interaction; Benchmark (surveying); Interface (matter); Modularity (biology); State (computer science); Theoretical computer science; Knowledge-based systems; Knowledge management","score_opus":0.018663867601655212,"score_gpt":0.2543226118598179,"score_spread":0.2356587442581627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122771415","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.074986845,0.00012845648,0.9206684,0.00038927505,0.000026064063,0.000045303546,0.00010174457,0.0017370526,0.0019168145],"genre_scores_gemma":[0.82526565,0.00024342036,0.17092061,0.00017767434,0.000032110856,0.00008839784,0.0003540278,0.00018041622,0.002737675],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994617,0.00012099553,0.000033859604,0.00023284095,0.000078949386,0.00007162652],"domain_scores_gemma":[0.99825567,0.0008874702,0.00018683549,0.0003822761,0.00016520945,0.00012260405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011111874,0.0010215573,0.00061264075,0.00052425347,0.0004392808,0.0015249961,0.0015199598,0.0011799128,0.0017627531],"category_scores_gemma":[0.005538229,0.0010097842,0.0011157511,0.00037435553,0.0015353763,0.005807673,0.0026227091,0.0020492647,0.0004850363],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022300206,0.00017548107,0.004079678,0.0001543559,0.00017442211,0.00023346668,0.0006777255,0.77221817,0.02231121,0.048365965,0.0012377189,0.1501488],"study_design_scores_gemma":[0.000007076161,0.0000421258,0.00041068238,0.000009144631,0.000026307256,0.0000257233,0.000024827123,0.96958274,0.0025584653,0.026792467,0.0005093741,0.0000111266645],"about_ca_topic_score_codex":0.008005107,"about_ca_topic_score_gemma":0.012351773,"teacher_disagreement_score":0.008005107,"about_ca_system_score_codex":0.0009993755,"about_ca_system_score_gemma":0.001042303,"threshold_uncertainty_score":0.015917063},"labels":[],"label_agreement":null},{"id":"W3124643392","doi":"","title":"Multiple Choice Question Answering using a Large Corpus of Information","year":2020,"lang":"en","type":"dissertation","venue":"University of Minnesota Digital Conservancy (University of Minnesota)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Ministère de la Santé et des Services sociaux; California Institute of Technology; National Aeronautics and Space Administration","keywords":"Computer science; Question answering; Artificial intelligence; Context (archaeology); Sentence; Knowledge base; Natural language; Embedding; Natural language processing; Information retrieval; Geography","score_opus":0.015227175305042214,"score_gpt":0.20077492691949847,"score_spread":0.18554775161445625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124643392","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5618542,0.015678467,0.15214038,0.009074159,0.0010942378,0.0018626441,0.20215029,0.018012367,0.038133256],"genre_scores_gemma":[0.47792003,0.0023335982,0.17515402,0.0010941986,0.0005860935,0.0017223981,0.3276554,0.0008121691,0.012722139],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99290943,0.003767284,0.00040022127,0.0017653571,0.0008920363,0.00026579638],"domain_scores_gemma":[0.95676625,0.036956232,0.00058310083,0.0022736795,0.002768274,0.0006525732],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005564447,0.0013160927,0.0017004319,0.004474325,0.0021116321,0.0035814203,0.001385007,0.0023985545,0.010399811],"category_scores_gemma":[0.030617753,0.00081332936,0.0009694374,0.0050272737,0.00064697483,0.0060191667,0.0027534203,0.002285298,0.009115929],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018456808,0.0019404858,0.022944443,0.0032036651,0.0005473569,0.0017664265,0.005914636,0.009747988,0.020528734,0.0055475524,0.37320304,0.55280995],"study_design_scores_gemma":[0.0011372863,0.0012444037,0.10783403,0.0011307034,0.0010618607,0.002053714,0.012450963,0.47134647,0.032643434,0.046053264,0.32248852,0.0005554305],"about_ca_topic_score_codex":0.008079968,"about_ca_topic_score_gemma":0.015473501,"teacher_disagreement_score":0.010399811,"about_ca_system_score_codex":0.0015092733,"about_ca_system_score_gemma":0.0018005179,"threshold_uncertainty_score":0.034790874},"labels":[],"label_agreement":null},{"id":"W3124933699","doi":"10.1016/j.ipm.2021.102503","title":"Learning to rank implicit entities on Twitter","year":2021,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Relevance (law); Graph; Representation (politics); Rank (graph theory); Learning to rank; Context (archaeology); Feature (linguistics); Natural language processing; Entity linking; Artificial intelligence; Knowledge base; Ranking (information retrieval); Theoretical computer science; Mathematics","score_opus":0.015156604224810487,"score_gpt":0.2522781744710504,"score_spread":0.2371215702462399,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124933699","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6224643,0.0048450846,0.34531602,0.005202845,0.0008080973,0.0002810091,0.008891652,0.0031876396,0.0090034725],"genre_scores_gemma":[0.9492143,0.00086926075,0.03353077,0.00018654957,0.00068476994,0.00009931384,0.008863958,0.00009062467,0.0064604725],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989661,0.00035124487,0.000081574435,0.00020516026,0.00021564106,0.00018024292],"domain_scores_gemma":[0.99706537,0.0017792552,0.00025721992,0.00026456473,0.00047050323,0.0001631779],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014346478,0.0009958252,0.0010012024,0.0036107926,0.00079868996,0.0017646479,0.0008769961,0.0011663912,0.0019934163],"category_scores_gemma":[0.005831614,0.00035849525,0.0007005115,0.0028966805,0.00039313515,0.0031356628,0.0011779565,0.0013522215,0.0022871634],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002059334,0.0010235538,0.13511114,0.000755922,0.0009097709,0.0004664368,0.0007118561,0.086114235,0.017577043,0.015402051,0.07260334,0.66726536],"study_design_scores_gemma":[0.000057208938,0.00019155898,0.009301562,0.000045235025,0.00014495556,0.00011291024,0.00029938115,0.96294004,0.004418539,0.01586423,0.006592695,0.000031641204],"about_ca_topic_score_codex":0.0059344354,"about_ca_topic_score_gemma":0.015894193,"teacher_disagreement_score":0.0059344354,"about_ca_system_score_codex":0.00063307385,"about_ca_system_score_gemma":0.00093386613,"threshold_uncertainty_score":0.011799812},"labels":[],"label_agreement":null},{"id":"W3125491423","doi":"10.1145/3442381.3449948","title":"Enquire One’s Parent and Child Before Decision: Fully Exploit Hierarchical Structure for Self-Supervised Taxonomy Expansion","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Taxonomy (biology); Computer science; Correctness; Artificial intelligence; Mean reciprocal rank; Exploit; Natural language processing; Information retrieval; Machine learning; Theoretical computer science; Algorithm","score_opus":0.04050101803570174,"score_gpt":0.2545640240924028,"score_spread":0.21406300605670106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3125491423","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16550772,0.004882856,0.8058458,0.0009943888,0.00021818485,0.0004398392,0.0022092052,0.011737079,0.008165041],"genre_scores_gemma":[0.5116209,0.00065217155,0.47098243,0.0006516468,0.00017938037,0.00025495308,0.008529738,0.000431536,0.0066973353],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99876875,0.00027373657,0.00007073359,0.0004210436,0.0003304631,0.00013518106],"domain_scores_gemma":[0.9977635,0.00088496925,0.0002178557,0.0004916079,0.0004821683,0.00015983552],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00151287,0.0010118614,0.00083249976,0.0024188843,0.0010537992,0.00071029866,0.0018976948,0.001191755,0.0016551827],"category_scores_gemma":[0.004291336,0.00036992136,0.0009382868,0.0021394864,0.0004973768,0.003815139,0.001489765,0.0012731272,0.0013886801],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021341407,0.00032788952,0.010415796,0.00028337125,0.00012192795,0.00023369645,0.0008376227,0.03102685,0.012514419,0.0065068286,0.033031303,0.90448684],"study_design_scores_gemma":[0.00006146377,0.00012329908,0.0041307583,0.000072778304,0.000091991365,0.00028818852,0.00032453256,0.95288795,0.0059281597,0.021066325,0.014986619,0.000037823527],"about_ca_topic_score_codex":0.010395867,"about_ca_topic_score_gemma":0.025024036,"teacher_disagreement_score":0.010395867,"about_ca_system_score_codex":0.0008612354,"about_ca_system_score_gemma":0.0018804086,"threshold_uncertainty_score":0.020670712},"labels":[],"label_agreement":null},{"id":"W3125846137","doi":"10.18280/ria.340606","title":"AQG: Arabic Question Generator","year":2020,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Generator (circuit theory); Natural language processing; Artificial intelligence; Arabic; Sentence; Grammar; Set (abstract data type); Process (computing); Semantics (computer science); Lexicon; Programming language; Linguistics","score_opus":0.06436756330711148,"score_gpt":0.2753353653220383,"score_spread":0.21096780201492682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3125846137","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010243783,0.0005023421,0.43021727,0.00070047355,0.0004061177,0.0013297376,0.012776572,0.53289354,0.010930187],"genre_scores_gemma":[0.15814838,0.00053170545,0.73902947,0.0011517089,0.00026489733,0.0021855389,0.05767047,0.018057588,0.022960262],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99867356,0.00047346996,0.00014545252,0.00035818637,0.0002812012,0.000068181194],"domain_scores_gemma":[0.99612683,0.0017387277,0.00019868846,0.0006651705,0.0010769845,0.0001935725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023221604,0.0019342435,0.0007624007,0.0024614118,0.00045672234,0.0012244385,0.0017585476,0.0012058774,0.045647226],"category_scores_gemma":[0.009654725,0.00056103343,0.00084411574,0.000978909,0.00050688855,0.0023154942,0.0022047993,0.0012294404,0.02999766],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093654933,0.00026771383,0.002846359,0.0018415713,0.000081003054,0.0005909796,0.0011650282,0.003337978,0.029668866,0.010844995,0.26817232,0.6802466],"study_design_scores_gemma":[0.0007940984,0.00065327226,0.0056037684,0.00036323932,0.00014045119,0.002481159,0.0006963345,0.22477923,0.112833254,0.044596,0.6068094,0.00024985056],"about_ca_topic_score_codex":0.0014213901,"about_ca_topic_score_gemma":0.0006005719,"teacher_disagreement_score":0.045647226,"about_ca_system_score_codex":0.0005886494,"about_ca_system_score_gemma":0.0007152406,"threshold_uncertainty_score":0.1527052},"labels":[],"label_agreement":null},{"id":"W3126195334","doi":"10.18653/v1/2021.eacl-main.224","title":"Unsupervised Abstractive Summarization of Bengali Text Documents","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Bengali; Automatic summarization; Computer science; Natural language processing; Unavailability; Artificial intelligence; Multi-document summarization; Information retrieval","score_opus":0.026213582583161227,"score_gpt":0.2699944682918565,"score_spread":0.2437808857086953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3126195334","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1212525,0.028443126,0.79878366,0.0029653527,0.002104752,0.00065814913,0.014410192,0.011769985,0.019612273],"genre_scores_gemma":[0.46759814,0.009227896,0.4242663,0.00033757917,0.0026327714,0.0004794576,0.047920484,0.0016128471,0.045924537],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985512,0.00037687688,0.00015018028,0.00037284684,0.0003747022,0.00017411135],"domain_scores_gemma":[0.9976031,0.00065338286,0.0002170406,0.00034111072,0.0010839019,0.000101515354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010133982,0.0015553687,0.0014257908,0.004342025,0.00090169924,0.0026138606,0.0013779136,0.0007230869,0.0032426298],"category_scores_gemma":[0.0037285925,0.00039398775,0.0010080237,0.0037982552,0.00043032854,0.0016648386,0.00125021,0.0010179669,0.0048264936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073336705,0.000138855,0.0014081157,0.000799284,0.00030086876,0.00039298492,0.000858065,0.008066485,0.057154015,0.0041716313,0.043107536,0.8828689],"study_design_scores_gemma":[0.0002445614,0.0008010605,0.02860579,0.00033736118,0.0016961676,0.0011390545,0.0034173804,0.52374035,0.16185778,0.028498365,0.24938336,0.00027877925],"about_ca_topic_score_codex":0.006808841,"about_ca_topic_score_gemma":0.01374331,"teacher_disagreement_score":0.006808841,"about_ca_system_score_codex":0.00076930656,"about_ca_system_score_gemma":0.0012983927,"threshold_uncertainty_score":0.01353842},"labels":[],"label_agreement":null},{"id":"W3126209695","doi":"10.2196/23587","title":"Novel Graph-Based Model With Biaffine Attention for Family History Extraction From Clinical Text: Modeling Study","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Harbin Institute of Technology","keywords":"Computer science; Relationship extraction; Family history; Information extraction; Natural language processing; Sentence; Information retrieval; Artificial intelligence; Genogram; ENCODE; Graph; Theoretical computer science; Psychology; Medicine","score_opus":0.09582043481915976,"score_gpt":0.34475927771473003,"score_spread":0.24893884289557028,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3126209695","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09274996,0.0015861498,0.89699435,0.0014483063,0.00012973059,0.00017425754,0.0016623024,0.001978013,0.003276872],"genre_scores_gemma":[0.8250604,0.001196889,0.158749,0.0005000386,0.00012261292,0.00030467563,0.0030945605,0.0002030038,0.010768822],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995846,0.00008293479,0.000021526967,0.00020814023,0.00005622791,0.000046555],"domain_scores_gemma":[0.9988366,0.0007811242,0.00010033713,0.00006385673,0.0001781296,0.00003993431],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00069244567,0.00094434683,0.0006070235,0.0014174579,0.0004204227,0.00071290345,0.0013375278,0.0012430879,0.0026401193],"category_scores_gemma":[0.0026232668,0.00042332246,0.0012354014,0.0011146335,0.00046213198,0.0015989998,0.00069696095,0.0011935931,0.0007594863],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035383884,0.00025998603,0.009155284,0.00023635545,0.00020110808,0.00058632856,0.00046626732,0.76553804,0.0064878166,0.016909018,0.007351306,0.19245471],"study_design_scores_gemma":[0.000004278999,0.000013849501,0.00044829803,0.000004149247,0.000017582837,0.000028746326,0.000008478691,0.99591357,0.0002989292,0.0027803197,0.00047680817,0.000005030611],"about_ca_topic_score_codex":0.04445068,"about_ca_topic_score_gemma":0.046710547,"teacher_disagreement_score":0.04445068,"about_ca_system_score_codex":0.0015692359,"about_ca_system_score_gemma":0.0012023109,"threshold_uncertainty_score":0.08838385},"labels":[],"label_agreement":null},{"id":"W3128005332","doi":"10.1037/rev0000265","title":"Disentangling contextual diversity: Communicative need as a lexical organizer.","year":2021,"lang":"en","type":"article","venue":"Psychological Review","topic":"Topic Modeling","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Context (archaeology); Lexicon; PsycINFO; Linguistics; Diversity (politics); Word lists by frequency; Psychology; Context effect; Lexical item; Cognitive psychology; Word (group theory); Computer science; Natural language processing; Sociology; History","score_opus":0.11862667975861994,"score_gpt":0.37841424539838875,"score_spread":0.2597875656397688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3128005332","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84646153,0.0015428636,0.11612197,0.0012165445,0.00006308148,0.00015890476,0.00054070214,0.00015146857,0.033742942],"genre_scores_gemma":[0.9882901,0.00010629186,0.010921885,0.00008309578,0.00002653164,0.000076939024,0.00014009347,0.000029720855,0.00032542032],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974953,0.0009393967,0.0001616006,0.0006436257,0.0006234463,0.00013662425],"domain_scores_gemma":[0.9829838,0.011164646,0.0023936604,0.0016513366,0.0009577349,0.0008487667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003081128,0.00068084267,0.0006908673,0.0043292693,0.0013530685,0.0039886567,0.0009061465,0.00085203594,0.002222568],"category_scores_gemma":[0.024590055,0.0005406054,0.0008552594,0.002616659,0.003921663,0.010570529,0.0057645217,0.0014693968,0.00021918632],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015317654,0.00045389315,0.2463736,0.0016821058,0.00085666234,0.0013036828,0.07994587,0.008100725,0.040233277,0.3047403,0.002178284,0.31259984],"study_design_scores_gemma":[0.00013501618,0.00070865033,0.28150985,0.00040595152,0.00079971296,0.001994275,0.024814617,0.061188616,0.008941592,0.6023827,0.016786158,0.00033281706],"about_ca_topic_score_codex":0.0021906882,"about_ca_topic_score_gemma":0.0025298896,"teacher_disagreement_score":0.0043292693,"about_ca_system_score_codex":0.0012392411,"about_ca_system_score_gemma":0.0008278042,"threshold_uncertainty_score":0.016294777},"labels":[],"label_agreement":null},{"id":"W3129137605","doi":"","title":"Modular Length Control for Sentence Generation.","year":2020,"lang":"en","type":"article","venue":"The European Symposium on Artificial Neural Networks","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of Guelph","funders":"","keywords":"Modular design; Computer science; Control (management); Programming language; Artificial intelligence","score_opus":0.04530803466431348,"score_gpt":0.23391705668300009,"score_spread":0.1886090220186866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3129137605","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021354526,0.0011452743,0.9623492,0.00035678453,0.00040986613,0.00015892394,0.0007293346,0.010091349,0.0034048578],"genre_scores_gemma":[0.54733896,0.00040476705,0.44100934,0.00030933923,0.0006073069,0.0005398492,0.0026288843,0.0016385681,0.0055229864],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99821883,0.0006439623,0.00013367504,0.00047186654,0.0003588024,0.00017291117],"domain_scores_gemma":[0.99253654,0.00416708,0.00042565545,0.0010046801,0.0015897054,0.00027629547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003320952,0.0010790614,0.00078006455,0.001164174,0.00064055226,0.0012878145,0.0019729722,0.001105417,0.008120217],"category_scores_gemma":[0.01669122,0.00043427185,0.00064617896,0.00091771956,0.0006074839,0.0024031687,0.0018376792,0.0018352702,0.003640827],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012498667,0.00025292812,0.0013072832,0.00032922378,0.000117106065,0.000088261266,0.00023009874,0.0533032,0.04053337,0.016400872,0.020655353,0.8655325],"study_design_scores_gemma":[0.00009474123,0.00020462665,0.0009310862,0.00003501417,0.000069556336,0.000059783273,0.000044630913,0.9368528,0.03210223,0.024063323,0.0055094715,0.000032770466],"about_ca_topic_score_codex":0.0026033327,"about_ca_topic_score_gemma":0.0043201833,"teacher_disagreement_score":0.008120217,"about_ca_system_score_codex":0.00092198333,"about_ca_system_score_gemma":0.0012111916,"threshold_uncertainty_score":0.027164817},"labels":[],"label_agreement":null},{"id":"W3129597052","doi":"10.2196/24678","title":"Extracting Drug Names and Associated Attributes From Discharge Summaries: Text Mining Study","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Engineering and Physical Sciences Research Council; King's College London; Saudi Arabian Cultural Bureau","keywords":"CRFS; Computer science; Conditional random field; Named-entity recognition; Natural language processing; Artificial intelligence; Word embedding; Context (archaeology); Biomedical text mining; Relationship extraction; Task (project management); Information extraction; Deep learning; Word (group theory); F1 score; Machine learning; Embedding; Text mining","score_opus":0.026157395346990336,"score_gpt":0.2886754367506368,"score_spread":0.26251804140364643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3129597052","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96527594,0.0015191279,0.016469436,0.0008440287,0.00007142166,0.0004477924,0.013715495,0.0006275048,0.0010292465],"genre_scores_gemma":[0.8913017,0.0015034662,0.069342054,0.0002612914,0.00015567896,0.00033680865,0.03581202,0.000065658176,0.0012213434],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997531,0.0005894658,0.0005870507,0.0006744442,0.0005245948,0.000093355935],"domain_scores_gemma":[0.97870535,0.014617788,0.00275827,0.0013758074,0.0021750382,0.00036768866],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003286394,0.00051722635,0.0005017854,0.0043612327,0.00046340792,0.00078898884,0.00089729147,0.00097179826,0.0008068844],"category_scores_gemma":[0.015336052,0.00018865224,0.0008959566,0.0033670445,0.00033276572,0.0014949384,0.00087139517,0.00070408545,0.0007510977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011616062,0.0027226135,0.58237565,0.0021116275,0.0003880771,0.0028405355,0.0020312571,0.010954225,0.008752165,0.0011618949,0.015464421,0.370036],"study_design_scores_gemma":[0.00033173,0.0016805497,0.581966,0.00064503105,0.00089023967,0.008392774,0.006290422,0.30960274,0.04641329,0.0033998643,0.04016777,0.00021953425],"about_ca_topic_score_codex":0.003548526,"about_ca_topic_score_gemma":0.004736715,"teacher_disagreement_score":0.0043612327,"about_ca_system_score_codex":0.0006490697,"about_ca_system_score_gemma":0.0009779899,"threshold_uncertainty_score":0.017380297},"labels":[],"label_agreement":null},{"id":"W3129849415","doi":"10.1007/978-3-030-67664-3_35","title":"AMQAN: Adaptive Multi-Attention Question-Answer Networks for Answer Selection","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Question answering; Automatic summarization; Selection (genetic algorithm); Subject (documents); SemEval; Questions and answers; Information retrieval; Artificial intelligence; Data science; World Wide Web; Task (project management)","score_opus":0.02582063519263548,"score_gpt":0.2638765607864926,"score_spread":0.2380559255938571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3129849415","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020167818,0.0022014284,0.93369126,0.0009855215,0.00060555735,0.0004394484,0.0047753267,0.030766705,0.006366926],"genre_scores_gemma":[0.25365517,0.0010434792,0.7028281,0.0011580698,0.000482435,0.0011533569,0.014347312,0.0015424027,0.02378961],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992786,0.00019028768,0.000035302724,0.00028044413,0.0001403108,0.00007507712],"domain_scores_gemma":[0.99823856,0.0011654623,0.000054190215,0.00018985919,0.0002492308,0.000102826234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016215807,0.0013594235,0.0013412008,0.0015751303,0.00085509766,0.0013731328,0.003037132,0.002159216,0.018683279],"category_scores_gemma":[0.0065341187,0.0007257703,0.0010583829,0.0014919058,0.00047443892,0.003144961,0.002817007,0.0024209253,0.0061397417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010579034,0.00040425945,0.0014909129,0.0004644514,0.00015925145,0.00013494925,0.00022163834,0.04449568,0.015430949,0.013860722,0.09558989,0.8266894],"study_design_scores_gemma":[0.00006407784,0.00006856635,0.00039517987,0.000024906942,0.000040059593,0.000042306117,0.00003235128,0.96076,0.0045117405,0.024351152,0.009692427,0.000017277003],"about_ca_topic_score_codex":0.008660194,"about_ca_topic_score_gemma":0.012681894,"teacher_disagreement_score":0.018683279,"about_ca_system_score_codex":0.001133947,"about_ca_system_score_gemma":0.0009786385,"threshold_uncertainty_score":0.06250179},"labels":[],"label_agreement":null},{"id":"W3130765241","doi":"10.48550/arxiv.2102.11417","title":"Parallelizing Legendre Memory Unit Training","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Leverage (statistics); Recurrent neural network; Machine translation; Artificial intelligence; Inference; Benchmark (surveying); Transformer; Invariant (physics); Machine learning; Artificial neural network; Mathematics","score_opus":0.20280706916611443,"score_gpt":0.2042012480033702,"score_spread":0.0013941788372557606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3130765241","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07368549,0.0009138156,0.89834815,0.000550234,0.00035461088,0.00010890306,0.000415814,0.01641578,0.009207224],"genre_scores_gemma":[0.60205835,0.0003593215,0.38123628,0.00043778098,0.000121287376,0.00022625305,0.0017235659,0.0010051122,0.012832003],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996069,0.00008242875,0.00002632828,0.00013209601,0.00008891644,0.00006341598],"domain_scores_gemma":[0.9994537,0.0001671782,0.000030674306,0.00017363964,0.00014987115,0.000024901698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00080359186,0.0008762605,0.0006607528,0.0005107645,0.00031854692,0.0007491473,0.0018454143,0.0007860943,0.0053577158],"category_scores_gemma":[0.002942004,0.00036497382,0.0005758162,0.00057758123,0.00044650302,0.0017516387,0.0009915857,0.0013053232,0.0022016156],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003357317,0.00014245007,0.00175961,0.00016783486,0.000119772194,0.00021804516,0.0001743197,0.31043485,0.02390558,0.016648548,0.018472776,0.6276205],"study_design_scores_gemma":[0.000012421416,0.000036742735,0.00016123739,0.0000071919403,0.000013157119,0.000034185017,0.000016936792,0.98376137,0.0099869585,0.0033045448,0.0026561338,0.00000913863],"about_ca_topic_score_codex":0.008576212,"about_ca_topic_score_gemma":0.016064977,"teacher_disagreement_score":0.008576212,"about_ca_system_score_codex":0.000792521,"about_ca_system_score_gemma":0.0011286906,"threshold_uncertainty_score":0.017923355},"labels":[],"label_agreement":null},{"id":"W3132704983","doi":"10.48550/arxiv.2102.12887","title":"Significant Improvements over the State of the Art? A Case Study of the MS MARCO Document Ranking Leaderboard","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); University of Waterloo","funders":"","keywords":"Ranking (information retrieval); Metric (unit); Context (archaeology); Computer science; Artificial intelligence; State (computer science); Information retrieval; Order (exchange); sort; Ask price; Machine learning; Learning to rank; Algorithm; Engineering; History; Operations management; Economics","score_opus":0.061682475121206934,"score_gpt":0.20014698999348118,"score_spread":0.13846451487227424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3132704983","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9330349,0.0068921354,0.01805722,0.008066981,0.00030990012,0.00019084364,0.002386566,0.001654959,0.029406592],"genre_scores_gemma":[0.96659863,0.00069150596,0.024995446,0.00045567288,0.00035958042,0.00008517886,0.002043581,0.00038320306,0.0043870956],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9901423,0.0060165594,0.00036699724,0.0008951431,0.002144893,0.0004342037],"domain_scores_gemma":[0.95806944,0.030510131,0.001954667,0.0040808213,0.0038005777,0.0015842887],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013778943,0.0007377259,0.0010363474,0.0024546047,0.0021983595,0.0034699026,0.0010684996,0.0013368317,0.0025978165],"category_scores_gemma":[0.044659253,0.00019040369,0.0005645582,0.0036495898,0.0016881597,0.0038022448,0.0014839465,0.0019989184,0.0012765123],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004704906,0.0020963587,0.10523332,0.002119992,0.00069980224,0.0027758419,0.010785821,0.052701645,0.014483702,0.045735255,0.13769266,0.6209707],"study_design_scores_gemma":[0.0010312751,0.0062785954,0.28146976,0.0007790406,0.0005302422,0.0028955957,0.020496983,0.34634867,0.03431969,0.07493278,0.23029268,0.00062459207],"about_ca_topic_score_codex":0.0069700833,"about_ca_topic_score_gemma":0.013817134,"teacher_disagreement_score":0.9862211,"about_ca_system_score_codex":0.0015920595,"about_ca_system_score_gemma":0.00087631476,"threshold_uncertainty_score":0.07287085},"labels":[],"label_agreement":null},{"id":"W3132891169","doi":"10.1007/978-3-662-45912-6_7","title":"Meaning–Focused and Quantum–Inspired Information Retrieval","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; Kelowna General Hospital; University of British Columbia","funders":"","keywords":"Computer science; Meaning (existential); Formalism (music); Scheme (mathematics); Natural language processing; Representation (politics); Artificial intelligence; Information retrieval; Identification (biology); Quantum; Epistemology; Mathematics; Philosophy; Quantum mechanics; Physics","score_opus":0.018553993762251614,"score_gpt":0.22672798980814246,"score_spread":0.20817399604589085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3132891169","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031037437,0.040200308,0.83660257,0.005384354,0.0013986991,0.00010876796,0.0002629948,0.0005631908,0.08444173],"genre_scores_gemma":[0.5623706,0.022571273,0.35421762,0.0016278224,0.003048223,0.00019009295,0.0009114017,0.00034415742,0.054718863],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931455,0.00021718063,0.00004788904,0.00011896261,0.00024836016,0.000053063566],"domain_scores_gemma":[0.9989549,0.0006060374,0.00006233374,0.00018829426,0.0001515102,0.000037012367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011719107,0.00041169184,0.00077826803,0.0013829,0.0006815523,0.0026320608,0.0013372857,0.001219299,0.005986374],"category_scores_gemma":[0.0039630304,0.00037010157,0.0007713499,0.0020334783,0.0016673958,0.0052618855,0.001512682,0.0014583159,0.0016283039],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006061744,0.000047656282,0.00020743809,0.0006857366,0.00005800098,0.000060784736,0.00046434542,0.0072968197,0.009658173,0.7809637,0.014199554,0.1862971],"study_design_scores_gemma":[0.000011675407,0.00003296193,0.00040598674,0.000064202926,0.000029433093,0.00018089754,0.00008493242,0.060606584,0.0033741428,0.9069198,0.028250402,0.000039049075],"about_ca_topic_score_codex":0.00044189088,"about_ca_topic_score_gemma":0.00051850063,"teacher_disagreement_score":0.005986374,"about_ca_system_score_codex":0.00094173214,"about_ca_system_score_gemma":0.00048171976,"threshold_uncertainty_score":0.020026386},"labels":[],"label_agreement":null},{"id":"W3133042814","doi":"10.2196/22797","title":"A Hybrid Model for Family History Information Identification and Relation Extraction: Development and Evaluation of an End-to-End Information Extraction System","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Cancer Institute","keywords":"Computer science; Information extraction; Relationship extraction; Artificial intelligence; Machine learning; Natural language processing; Context (archaeology); Heuristics; Named-entity recognition; Task (project management); Identification (biology); Information retrieval","score_opus":0.054910600379122286,"score_gpt":0.3097503389832227,"score_spread":0.2548397386041004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133042814","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1262374,0.0010376406,0.8098772,0.0016207739,0.00033916422,0.001016598,0.0067979526,0.048971817,0.0041014673],"genre_scores_gemma":[0.32772028,0.00042405896,0.64626974,0.0012331908,0.00010092591,0.0010131391,0.01239613,0.00060848263,0.010234075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99907386,0.00018617083,0.00008980492,0.00037959527,0.0002013,0.00006922057],"domain_scores_gemma":[0.9982134,0.0009943444,0.00007119486,0.00020071829,0.00043057368,0.00008980993],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016271825,0.0011045334,0.0009338831,0.0009791049,0.0006566966,0.0012196033,0.0022627036,0.0015908093,0.0046556843],"category_scores_gemma":[0.0038628865,0.00043014952,0.0010901432,0.0007080835,0.00033078852,0.0020263146,0.0012994708,0.0013488701,0.0028196818],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015968948,0.0014228752,0.01446451,0.0004615406,0.00052277104,0.0015653464,0.00044960124,0.17309643,0.02187483,0.0042124884,0.03884722,0.7414854],"study_design_scores_gemma":[0.000054378084,0.000101790356,0.0010717958,0.000017552966,0.0000655751,0.00013293473,0.000032457632,0.9886818,0.0052068494,0.0018101119,0.0028026528,0.000022082086],"about_ca_topic_score_codex":0.02113808,"about_ca_topic_score_gemma":0.02512093,"teacher_disagreement_score":0.02113808,"about_ca_system_score_codex":0.0011918542,"about_ca_system_score_gemma":0.0017428567,"threshold_uncertainty_score":0.042030096},"labels":[],"label_agreement":null},{"id":"W3133607072","doi":"10.1111/ijcp.14140","title":"Plain‐language summaries: An essential component to promote knowledge translation","year":2021,"lang":"en","type":"article","venue":"International Journal of Clinical Practice","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Chiropractic Association; University of Manitoba; Manitoba Health","funders":"","keywords":"Publishing; Work (physics); Knowledge translation; Medicine; Public relations; Plain language; Sociology of scientific knowledge; Scientific literature; Engineering ethics; Knowledge management; Sociology; Social science; Political science; Computer science","score_opus":0.12063278184568481,"score_gpt":0.4787449711156056,"score_spread":0.3581121892699208,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133607072","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0118683735,0.022176981,0.46003175,0.32255074,0.048590932,0.015681874,0.010960567,0.017154917,0.090983875],"genre_scores_gemma":[0.07900353,0.025859544,0.7540623,0.04896316,0.026421243,0.019848287,0.008801927,0.008719997,0.028320001],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.80607,0.13505743,0.030595712,0.0041713323,0.022788811,0.0013166261],"domain_scores_gemma":[0.31590042,0.4535583,0.038308695,0.07716773,0.10711641,0.007948527],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.18438499,0.0018313559,0.0028495004,0.011701146,0.003315374,0.018474763,0.003692132,0.0066203927,0.037166744],"category_scores_gemma":[0.6010312,0.0022532053,0.0016004964,0.0086490745,0.00413354,0.023749141,0.012048652,0.012515325,0.05022983],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046580512,0.00039000897,0.0007450977,0.011270026,0.00014744686,0.00052233634,0.029569928,0.00044881328,0.002988015,0.03152355,0.39045757,0.5314713],"study_design_scores_gemma":[0.00018782503,0.00023812811,0.0012521626,0.011678087,0.000108680004,0.0006074734,0.006122592,0.0008534025,0.0019402449,0.042378616,0.93443906,0.00019369688],"about_ca_topic_score_codex":0.0009106753,"about_ca_topic_score_gemma":0.0013689297,"teacher_disagreement_score":0.815615,"about_ca_system_score_codex":0.0031044923,"about_ca_system_score_gemma":0.01726735,"threshold_uncertainty_score":0.97513264},"labels":[],"label_agreement":null},{"id":"W3133879190","doi":"10.48550/arxiv.2103.02190","title":"An Iterative Contextualization Algorithm with Second-Order Attention","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Contextualization; Computer science; Sentence; Value (mathematics); Embedding; Encoding (memory); Context (archaeology); Order (exchange); Algorithm; Artificial intelligence; Theoretical computer science; Natural language processing; Machine learning","score_opus":0.05022255834153722,"score_gpt":0.19474064017306572,"score_spread":0.1445180818315285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133879190","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007672217,0.00021918135,0.9871783,0.00018419571,0.00006898768,0.00013511823,0.00006523331,0.0030207117,0.001456047],"genre_scores_gemma":[0.15103683,0.00018519316,0.84142566,0.0002895181,0.00016665152,0.00041871658,0.00055913895,0.00055435766,0.0053639878],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987602,0.00029998124,0.000068730435,0.00046281575,0.0002501251,0.00015816046],"domain_scores_gemma":[0.99874735,0.00050920923,0.00008218894,0.00027570614,0.00032135277,0.00006420683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012528051,0.0016358716,0.001344754,0.0019885188,0.0008852033,0.0012543178,0.0020743418,0.0013588804,0.0066313655],"category_scores_gemma":[0.004438023,0.0007454977,0.0011823192,0.0019115319,0.0007560941,0.0019919712,0.002561,0.0017053953,0.0022548446],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024680357,0.00016697719,0.0013899786,0.00017770307,0.00013572042,0.00012183781,0.00049540587,0.04606727,0.01843865,0.026943235,0.011221219,0.89459527],"study_design_scores_gemma":[0.000072233765,0.00014599695,0.0008243661,0.000028324292,0.000087979555,0.00015156456,0.00012126911,0.93250406,0.010542064,0.044851243,0.010632004,0.000038926337],"about_ca_topic_score_codex":0.006681126,"about_ca_topic_score_gemma":0.013044196,"teacher_disagreement_score":0.006681126,"about_ca_system_score_codex":0.0010268708,"about_ca_system_score_gemma":0.0020305328,"threshold_uncertainty_score":0.022184134},"labels":[],"label_agreement":null},{"id":"W3134614468","doi":"10.31234/osf.io/gvr6m","title":"Do We Need Neural Models to Explain Human Judgments of Acceptability?","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Social Science Fund of China; Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Sentence; Computer science; Natural language processing; Language model; Word (group theory); Artificial intelligence; Simple (philosophy); Variance (accounting); Similarity (geometry); n-gram; Sentence processing; Linguistics","score_opus":0.1065612792918263,"score_gpt":0.29397034668807787,"score_spread":0.18740906739625157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134614468","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6334606,0.0030043824,0.32507014,0.011640237,0.00041015525,0.000097330005,0.0016123882,0.0014626732,0.02324212],"genre_scores_gemma":[0.98103803,0.00039935348,0.015257457,0.00059993984,0.000099237484,0.000050593935,0.0004999956,0.0000928564,0.001962527],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999423,0.000271606,0.000015524402,0.00018257454,0.000056088527,0.000051238072],"domain_scores_gemma":[0.9962567,0.002240482,0.0003594147,0.00047055134,0.0004945753,0.00017829686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016650843,0.00069289125,0.00053383585,0.00053671515,0.00024344375,0.0018585585,0.000773473,0.00097261573,0.0034835285],"category_scores_gemma":[0.014833341,0.00045288017,0.00050285796,0.00034191363,0.00076147425,0.00506696,0.000515374,0.0018731492,0.0015603267],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088049483,0.000601396,0.17767476,0.0011165678,0.0013666522,0.0005467857,0.005379742,0.21075675,0.039086226,0.07905613,0.025309507,0.458225],"study_design_scores_gemma":[0.000038269398,0.00012870833,0.04534838,0.00011509109,0.000070775444,0.00020662924,0.0008113659,0.7666933,0.0023341894,0.17963646,0.004531523,0.00008530133],"about_ca_topic_score_codex":0.005408975,"about_ca_topic_score_gemma":0.0054352037,"teacher_disagreement_score":0.005408975,"about_ca_system_score_codex":0.00053949247,"about_ca_system_score_gemma":0.00043349818,"threshold_uncertainty_score":0.011653602},"labels":[],"label_agreement":null},{"id":"W3134665270","doi":"10.1145/3437963.3441667","title":"Pretrained Transformers for Text Ranking: BERT and Beyond","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":164,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Transformer; Computer science; Artificial intelligence; Ranking (information retrieval); Natural language processing; Question answering; Language model; Information retrieval; Machine learning; Natural language; Engineering","score_opus":0.018037171180899265,"score_gpt":0.24312119392371964,"score_spread":0.22508402274282038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134665270","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008559169,0.0036418175,0.9782532,0.0015134402,0.00017439843,0.0000712698,0.000244224,0.0026722099,0.0048704282],"genre_scores_gemma":[0.43245822,0.010714678,0.5252085,0.0013372577,0.0010061312,0.00029983732,0.0019088353,0.0012767324,0.025789829],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993948,0.00022602435,0.000038890623,0.000120679404,0.00015852573,0.00006106253],"domain_scores_gemma":[0.9977437,0.0012989774,0.000115917785,0.00042035384,0.00033454038,0.00008641028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015684584,0.0010932502,0.00091751537,0.0009416434,0.00039486008,0.0019175119,0.0013733526,0.0009995602,0.004829126],"category_scores_gemma":[0.0075276326,0.00048879994,0.0005765911,0.0014411083,0.0010812697,0.0064921067,0.0012482449,0.003244196,0.002751776],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002479355,0.00013250354,0.00086842314,0.00038425985,0.000075730226,0.00010269816,0.00025813785,0.10444113,0.009080168,0.14321154,0.019800255,0.72139734],"study_design_scores_gemma":[0.00002040326,0.00012589742,0.0004172041,0.00007113482,0.000031437303,0.000095524476,0.000065710505,0.7988894,0.006053129,0.18041809,0.013780807,0.00003122498],"about_ca_topic_score_codex":0.0036180601,"about_ca_topic_score_gemma":0.0051664663,"teacher_disagreement_score":0.004829126,"about_ca_system_score_codex":0.000998322,"about_ca_system_score_gemma":0.001076614,"threshold_uncertainty_score":0.016155064},"labels":[],"label_agreement":null},{"id":"W3134942116","doi":"10.1109/icsc50631.2021.00056","title":"Using Conditional Sentence Representation in Pointer Networks for Sentence Ordering","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Sentence; Computer science; Pointer (user interface); Natural language processing; Artificial intelligence; Representation (politics)","score_opus":0.07775070764176942,"score_gpt":0.32556789947407566,"score_spread":0.24781719183230624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134942116","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039980546,0.0013183559,0.9442857,0.0009148276,0.00024019464,0.00026927656,0.0030439948,0.005603742,0.004343432],"genre_scores_gemma":[0.6023305,0.0014576284,0.3708143,0.00067525037,0.0004250076,0.0009594689,0.013587751,0.00051679957,0.009233346],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993994,0.00020596792,0.000049220936,0.00018924686,0.00010293776,0.000053268155],"domain_scores_gemma":[0.9981025,0.0010811121,0.00019260864,0.00018031188,0.00034639164,0.00009708088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012044096,0.0012806986,0.00063597603,0.00267079,0.00064205704,0.0010115481,0.001557141,0.0011888818,0.005815045],"category_scores_gemma":[0.0064309295,0.00040949386,0.0010538845,0.002348948,0.0005362401,0.004028722,0.0012039256,0.001854743,0.0016595536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008158169,0.00044821986,0.0052309493,0.0006766491,0.00020882879,0.00052637316,0.0009756133,0.19151212,0.022251403,0.06926969,0.028647771,0.6794366],"study_design_scores_gemma":[0.00003328424,0.00013144502,0.001005859,0.000057278372,0.00009496033,0.00011560601,0.00007343054,0.9342404,0.004032638,0.053579893,0.0065955436,0.000039564446],"about_ca_topic_score_codex":0.0064198575,"about_ca_topic_score_gemma":0.011711991,"teacher_disagreement_score":0.0064198575,"about_ca_system_score_codex":0.0013229658,"about_ca_system_score_gemma":0.0015336922,"threshold_uncertainty_score":0.019453287},"labels":[],"label_agreement":null},{"id":"W3135857378","doi":"10.2196/24020","title":"Extracting Family History Information From Electronic Health Records: Natural Language Processing Analysis","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Commonwealth Scientific and Industrial Research Organisation","keywords":"Natural history; Health records; Text messaging; Family history; Computer science; Natural language processing; Medicine; Data science; World Wide Web; Health care","score_opus":0.014316315972030631,"score_gpt":0.2881194847011004,"score_spread":0.27380316872906973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3135857378","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23354532,0.0038397096,0.72360384,0.00446973,0.00028714054,0.0011621898,0.019543441,0.010414013,0.0031344937],"genre_scores_gemma":[0.31032544,0.0013842436,0.6552504,0.0006944468,0.0002185003,0.0004974215,0.030027537,0.00016333887,0.0014386212],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974915,0.000967213,0.00035134068,0.00063286937,0.00046496984,0.00009216509],"domain_scores_gemma":[0.9861572,0.010949456,0.00086185016,0.00065173197,0.0012502369,0.00012960378],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036414391,0.0009200747,0.00049511006,0.0035941966,0.00053768366,0.001004147,0.0011274768,0.0008056854,0.0012823011],"category_scores_gemma":[0.012191767,0.00030278397,0.0010008418,0.0017452152,0.00039550828,0.0016104492,0.0010322797,0.0010851285,0.0011581839],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045917573,0.0005161692,0.03885076,0.0018897706,0.00025367527,0.002303447,0.0016675849,0.024061637,0.054034546,0.0017831126,0.02460551,0.8495746],"study_design_scores_gemma":[0.00015563483,0.0004984866,0.08885871,0.0006933829,0.00057501777,0.0036696417,0.0030707458,0.70446247,0.11856751,0.021373633,0.057830814,0.00024401925],"about_ca_topic_score_codex":0.0058468496,"about_ca_topic_score_gemma":0.008778275,"teacher_disagreement_score":0.0058468496,"about_ca_system_score_codex":0.00079569843,"about_ca_system_score_gemma":0.0016931277,"threshold_uncertainty_score":0.019258022},"labels":[],"label_agreement":null},{"id":"W3136109765","doi":"10.1145/3458553.3458563","title":"The neural hype, justified!","year":2019,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Eternity; Artificial intelligence; Epistemology; Philosophy","score_opus":0.01663961322052802,"score_gpt":0.2390605215602842,"score_spread":0.2224209083397562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136109765","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026811296,0.017029442,0.026125006,0.7847701,0.036146216,0.00002488172,0.0005860374,0.0010840204,0.13155329],"genre_scores_gemma":[0.22420648,0.0178856,0.025875362,0.40099573,0.022431588,0.00020248404,0.000921174,0.0035597398,0.3039218],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9962225,0.0017654051,0.0001548073,0.00055282377,0.0010233085,0.00028107016],"domain_scores_gemma":[0.9912361,0.0045928136,0.00045079514,0.0013474304,0.0016725273,0.0007004066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041518826,0.00085875124,0.00049981894,0.0008757877,0.0041780695,0.0069549214,0.0009274239,0.004651742,0.04072257],"category_scores_gemma":[0.03873734,0.00046678865,0.0003528067,0.00071962515,0.012063856,0.015831312,0.0038596236,0.014621171,0.023017328],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010372125,0.000028501512,0.00039553057,0.0001469271,0.000024734856,0.00014979497,0.0012131218,0.00012495254,0.0004983276,0.21003117,0.74567485,0.04160837],"study_design_scores_gemma":[0.000023392635,0.000036033704,0.00035662792,0.00030262722,0.00001158111,0.00034241666,0.0011773285,0.00068413356,0.00071040087,0.15599896,0.8403167,0.00003972313],"about_ca_topic_score_codex":0.0038440516,"about_ca_topic_score_gemma":0.005748328,"teacher_disagreement_score":0.04072257,"about_ca_system_score_codex":0.0017767146,"about_ca_system_score_gemma":0.0015006643,"threshold_uncertainty_score":0.13623059},"labels":[],"label_agreement":null},{"id":"W3136275640","doi":"10.1109/bigdata50022.2020.9378201","title":"Customizing Contextualized Language Models for Legal Document Reviews","year":2020,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Natural language processing; Domain (mathematical analysis); Artificial intelligence; Inference; Focus (optics); Task (project management); Security token; Language model; Question answering","score_opus":0.0689026083376294,"score_gpt":0.31200194988777435,"score_spread":0.24309934155014495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136275640","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12497465,0.0026379363,0.8303411,0.0009384929,0.0004900122,0.00051001966,0.00229738,0.033481352,0.0043289363],"genre_scores_gemma":[0.57087195,0.0010700551,0.41274726,0.00049972336,0.00038407018,0.0007065887,0.007989022,0.0018810444,0.0038502112],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99814403,0.00080845837,0.00013611169,0.00060782826,0.00017707502,0.00012647285],"domain_scores_gemma":[0.9963504,0.001819087,0.00024425102,0.0006098951,0.00082341145,0.000152921],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026461154,0.0012131776,0.0009151546,0.0015381991,0.00042950912,0.0014844821,0.0014292568,0.001077421,0.0025569957],"category_scores_gemma":[0.009445618,0.00065433455,0.0012566934,0.00095471035,0.00040315877,0.0025912751,0.0012210377,0.0015261314,0.0032466208],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092528766,0.00062999746,0.007564744,0.001018134,0.00069202733,0.000502992,0.0010404291,0.2583417,0.059194062,0.0068986965,0.033420518,0.6297714],"study_design_scores_gemma":[0.00004625674,0.00010828963,0.0010747068,0.000033449698,0.00010632685,0.00013041883,0.00014777078,0.9721594,0.013491432,0.0037293176,0.008922392,0.00005018364],"about_ca_topic_score_codex":0.006303597,"about_ca_topic_score_gemma":0.013146135,"teacher_disagreement_score":0.006303597,"about_ca_system_score_codex":0.0007760955,"about_ca_system_score_gemma":0.0017876525,"threshold_uncertainty_score":0.013994157},"labels":[],"label_agreement":null},{"id":"W3136363192","doi":"10.1016/j.csl.2022.101429","title":"On the effect of dropping layers of pre-trained transformer models","year":2022,"lang":"en","type":"article","venue":"Computer Speech & Language","topic":"Topic Modeling","field":"Computer Science","cited_by":91,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Transformer; Computer science; Limiting; Sentence; Artificial intelligence; Paraphrase; Machine learning; Natural language processing; Engineering","score_opus":0.010209370206985395,"score_gpt":0.2361752071494904,"score_spread":0.225965836942505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136363192","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55225635,0.010741075,0.38277832,0.004456219,0.0038667629,0.00035228164,0.0022250758,0.01836221,0.024961716],"genre_scores_gemma":[0.88128775,0.0015847647,0.09477663,0.0016504881,0.00023467094,0.000078275916,0.0027757536,0.002203322,0.015408248],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835545,0.00054733304,0.00010605688,0.00037824063,0.00026848505,0.000344414],"domain_scores_gemma":[0.9882004,0.008810316,0.00022343658,0.0012505098,0.0011667116,0.0003485485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040103067,0.0025430154,0.0012483858,0.0006874642,0.00091651717,0.0023163348,0.0016003912,0.0029419817,0.011600522],"category_scores_gemma":[0.028775083,0.0008883465,0.0008605853,0.0007098003,0.00096132193,0.004007506,0.0017392121,0.004294732,0.0024016693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009604643,0.0012805853,0.0073509645,0.00089980476,0.001031077,0.00079444004,0.00053044397,0.35849378,0.09318565,0.006643346,0.027110022,0.49307522],"study_design_scores_gemma":[0.0002535243,0.0006951978,0.0039884024,0.00016890018,0.00064573524,0.00025240157,0.0003313408,0.92283076,0.06118226,0.0044418103,0.005149716,0.000060039536],"about_ca_topic_score_codex":0.032517157,"about_ca_topic_score_gemma":0.06206566,"teacher_disagreement_score":0.032517157,"about_ca_system_score_codex":0.0011184601,"about_ca_system_score_gemma":0.002189918,"threshold_uncertainty_score":0.06465578},"labels":[],"label_agreement":null},{"id":"W3136687158","doi":"10.1162/tacl","title":"MasakhaNER: Named entity recognition for African languages","year":2021,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Named-entity recognition; Computer science; Natural language processing; Linguistics; Entity linking; Artificial intelligence; Philosophy; Engineering; Task (project management)","score_opus":0.0201344719929059,"score_gpt":0.2364009823153819,"score_spread":0.21626651032247599,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136687158","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045403384,0.0024359792,0.3964216,0.001798718,0.0008503368,0.0007869754,0.12694357,0.41380048,0.0115589695],"genre_scores_gemma":[0.20550863,0.0014244755,0.46909755,0.00069623784,0.00022941582,0.0011926121,0.28370306,0.012556378,0.025591735],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9991685,0.0001741562,0.00009340651,0.00028523273,0.0001704686,0.00010823375],"domain_scores_gemma":[0.99901307,0.000367034,0.00008103464,0.00026408737,0.00016925398,0.00010559954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014208514,0.0013607376,0.001130276,0.0020696684,0.0010180871,0.0021110745,0.001503414,0.0010456379,0.023535086],"category_scores_gemma":[0.003542874,0.0007511034,0.0013022244,0.002103394,0.00030709946,0.0048896126,0.0026890927,0.0014145755,0.019438857],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002119582,0.0002970383,0.0062481686,0.0014929572,0.00038155643,0.0008769887,0.0007347293,0.0046391394,0.03559159,0.012406191,0.44188374,0.49332836],"study_design_scores_gemma":[0.00057774305,0.0005520554,0.013687344,0.0003211018,0.00047770992,0.0015894846,0.0011319785,0.24178177,0.12952666,0.028965859,0.58109146,0.00029678806],"about_ca_topic_score_codex":0.0040782173,"about_ca_topic_score_gemma":0.0042197676,"teacher_disagreement_score":0.023535086,"about_ca_system_score_codex":0.0005751435,"about_ca_system_score_gemma":0.0010368313,"threshold_uncertainty_score":0.07873267},"labels":[],"label_agreement":null},{"id":"W3137721922","doi":"10.1609/aaai.v35i17.17829","title":"Deep Discourse Analysis for Generating Personalized Feedback in Intelligent Tutor Systems","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; TUTOR; Artificial intelligence; Classifier (UML); Graph; Intelligent tutoring system; Process (computing); Personalized learning; Natural language processing; Human–computer interaction; Teaching method; Mathematics education; Cooperative learning; Theoretical computer science; Psychology","score_opus":0.10135030651859706,"score_gpt":0.323220521530623,"score_spread":0.22187021501202592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3137721922","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.089329265,0.00069032685,0.90168923,0.0006923892,0.000051690542,0.00013791613,0.00030383925,0.005125215,0.0019800626],"genre_scores_gemma":[0.6902216,0.0002468706,0.305762,0.00015295952,0.000052480995,0.00018003212,0.0006423116,0.00025087176,0.0024908641],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979838,0.0010776477,0.00011620305,0.0003762316,0.00032366384,0.00012230605],"domain_scores_gemma":[0.99502003,0.0035275137,0.00033309145,0.00025719684,0.00074270077,0.00011947482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024880914,0.0009839807,0.00069765234,0.0020021114,0.0006732739,0.0012610796,0.0012879702,0.0013474912,0.0020487334],"category_scores_gemma":[0.010019339,0.0004306084,0.0005276275,0.0009417807,0.00069892575,0.0024572648,0.0014001161,0.0012484674,0.0006245289],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046739413,0.00038978274,0.0037015183,0.00051051803,0.00009674325,0.00030573164,0.0026741961,0.25914684,0.03619902,0.012857224,0.0050220215,0.678629],"study_design_scores_gemma":[0.000016644146,0.000049379636,0.0003746889,0.00001516207,0.000015871265,0.000016101272,0.00014356109,0.97726405,0.009415928,0.011367942,0.0013103046,0.000010355461],"about_ca_topic_score_codex":0.004821803,"about_ca_topic_score_gemma":0.0064787786,"teacher_disagreement_score":0.004821803,"about_ca_system_score_codex":0.0016542107,"about_ca_system_score_gemma":0.0011224417,"threshold_uncertainty_score":0.013158441},"labels":[],"label_agreement":null},{"id":"W3138967041","doi":"10.18653/v1/2021.naacl-main.43","title":"Open Domain Question Answering over Tables via Dense Retrieval","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Labrador Retriever; Information retrieval; Computer science; Context (archaeology); Domain (mathematical analysis); Precision and recall; Table (database); Question answering; Data mining; Mathematics; Geography; Medicine","score_opus":0.028085760024425074,"score_gpt":0.29289713553080904,"score_spread":0.26481137550638395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3138967041","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06692385,0.006565018,0.8774198,0.0036945206,0.00031596705,0.0004973116,0.013814956,0.019755654,0.011012947],"genre_scores_gemma":[0.47928083,0.0025958582,0.45792636,0.0009769913,0.0006200943,0.00038157642,0.047206584,0.00088703097,0.010124746],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99618894,0.0013714766,0.00031521314,0.0009086194,0.00085960457,0.00035622262],"domain_scores_gemma":[0.98993534,0.006198213,0.0002946218,0.0023761685,0.00091170974,0.00028387082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002828617,0.0012286408,0.0028303124,0.0053333333,0.0013712635,0.0052255834,0.0025789493,0.0019847788,0.010645146],"category_scores_gemma":[0.016103933,0.0011425072,0.0018092488,0.0066168006,0.0012296067,0.015813848,0.007117906,0.0023383775,0.0071321214],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012677301,0.0008364239,0.0040593217,0.0013631139,0.00042474133,0.00047745294,0.001677467,0.044309974,0.014213029,0.09904152,0.14994061,0.68238854],"study_design_scores_gemma":[0.00022245193,0.00019032147,0.0013653187,0.00012512539,0.00021092419,0.00029839808,0.0009505177,0.48542285,0.0076280483,0.4762982,0.027209757,0.00007807285],"about_ca_topic_score_codex":0.0073847314,"about_ca_topic_score_gemma":0.012083929,"teacher_disagreement_score":0.010645146,"about_ca_system_score_codex":0.0013824288,"about_ca_system_score_gemma":0.0016746679,"threshold_uncertainty_score":0.03561157},"labels":[],"label_agreement":null},{"id":"W3140853032","doi":"10.1007/978-3-030-72240-1_11","title":"Comparing Score Aggregation Approaches for Document Retrieval with Pretrained Transformers","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Transformer; Artificial intelligence; Question answering; Information retrieval; Machine learning","score_opus":0.05235121322312605,"score_gpt":0.24538324773872552,"score_spread":0.19303203451559947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3140853032","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2492177,0.008249388,0.72062534,0.00070242205,0.0004418943,0.00037925044,0.0013440421,0.010784286,0.008255657],"genre_scores_gemma":[0.78181374,0.0022498867,0.20448603,0.00013925298,0.00032723031,0.00018899936,0.0039744494,0.00075986446,0.0060606063],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99678826,0.0011335324,0.00031700646,0.0003997878,0.0010064938,0.00035493053],"domain_scores_gemma":[0.98953944,0.007093222,0.00023260861,0.0015871585,0.0013675875,0.00017995783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060982816,0.0012009386,0.0019031396,0.002884908,0.0006142278,0.0025065597,0.001817707,0.0013450794,0.003683409],"category_scores_gemma":[0.015433875,0.00048984704,0.0012056734,0.0034116793,0.0007083484,0.0058530173,0.0019197654,0.001479995,0.0015723476],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003186555,0.00052104704,0.0033275802,0.00048744457,0.00057375774,0.000088711524,0.00023687718,0.11618608,0.0072459355,0.014475814,0.012765162,0.840905],"study_design_scores_gemma":[0.00016237033,0.0005704115,0.0019697521,0.000026151316,0.00021716738,0.00009496183,0.00012556535,0.9717405,0.006157625,0.017261509,0.0016372093,0.000036697355],"about_ca_topic_score_codex":0.007844736,"about_ca_topic_score_gemma":0.007132871,"teacher_disagreement_score":0.007844736,"about_ca_system_score_codex":0.0012170566,"about_ca_system_score_gemma":0.0015854132,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3141030693","doi":"10.1109/wetice.2007.4407137","title":"The Roots and the Rationale behind the ALM Based Collaboration","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science","score_opus":0.01240663855113928,"score_gpt":0.2418082583584693,"score_spread":0.22940161980733004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3141030693","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017574282,0.00550733,0.42579564,0.09535059,0.0028862853,0.000466501,0.000266695,0.0009936591,0.45115897],"genre_scores_gemma":[0.62636435,0.0031613747,0.20363373,0.006789701,0.0019867343,0.0012985896,0.00025478625,0.00057520776,0.15593557],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9790733,0.013514804,0.00071538624,0.0018058893,0.0038618373,0.0010288176],"domain_scores_gemma":[0.9812897,0.010469115,0.0016422626,0.0028456633,0.002421058,0.0013321077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015322894,0.0006297685,0.00068616227,0.0024763567,0.0069786403,0.013663652,0.0029349267,0.00545547,0.021035442],"category_scores_gemma":[0.029202493,0.000785517,0.00078539684,0.0031374192,0.015932081,0.018075762,0.011643426,0.005823558,0.0072603677],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000021775691,0.000020080344,0.00022331419,0.00010454385,0.000003882231,0.000106613326,0.0037320706,0.00025709384,0.0002120261,0.971245,0.005447068,0.018626468],"study_design_scores_gemma":[0.0000306374,0.000034175428,0.00035654736,0.00033777708,0.00000984872,0.00033033406,0.0038627058,0.0038301414,0.0006275949,0.8060912,0.18445025,0.00003875524],"about_ca_topic_score_codex":0.0016958443,"about_ca_topic_score_gemma":0.0024465558,"teacher_disagreement_score":0.021035442,"about_ca_system_score_codex":0.0063872407,"about_ca_system_score_gemma":0.0048882836,"threshold_uncertainty_score":0.08103615},"labels":[],"label_agreement":null},{"id":"W3149251347","doi":"10.22215/etd/2020-14211","title":"Sequence Modeling with Linear Complexity","year":2020,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Automatic summarization; Kernel (algebra); Sequence (biology); Benchmark (surveying); Autoregressive model; Time complexity; Computational complexity theory; Task (project management); Artificial intelligence; Memory footprint; Quadratic equation; Machine translation; Algorithm; Machine learning; Theoretical computer science; Mathematics","score_opus":0.12940945773198947,"score_gpt":0.311235620889329,"score_spread":0.18182616315733952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3149251347","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003068174,0.0003505835,0.99306047,0.0003047183,0.000068722286,0.000048604925,0.00024882553,0.0009535625,0.0018964024],"genre_scores_gemma":[0.3197963,0.002100981,0.6472645,0.00058584363,0.0007640714,0.0010511825,0.0030408795,0.0010134752,0.024382787],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.998621,0.00036837245,0.00009315131,0.00035836152,0.00043725208,0.000121954115],"domain_scores_gemma":[0.9956801,0.002917852,0.00027407074,0.0005898148,0.00045585568,0.00008225165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012546203,0.0012611389,0.0012227218,0.0010370619,0.0005828789,0.0019056213,0.0016153443,0.0013097622,0.00792749],"category_scores_gemma":[0.009070303,0.0007349135,0.0018180695,0.0013249121,0.0008808135,0.0030643225,0.0015950669,0.0026003083,0.0036645075],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014241759,0.00006967533,0.0008856501,0.0002947653,0.00011195842,0.00022160265,0.00015972069,0.7192699,0.0035418186,0.1581433,0.00832965,0.108829506],"study_design_scores_gemma":[0.0000043556906,0.000008637495,0.000052503216,0.000004859879,0.0000061907963,0.000020752132,0.000005386234,0.9700398,0.00039774785,0.028006451,0.0014488029,0.0000046763575],"about_ca_topic_score_codex":0.008046319,"about_ca_topic_score_gemma":0.007450929,"teacher_disagreement_score":0.008046319,"about_ca_system_score_codex":0.0016878294,"about_ca_system_score_gemma":0.0014058294,"threshold_uncertainty_score":0.026520073},"labels":[],"label_agreement":null},{"id":"W3150524533","doi":"10.1007/978-3-030-73197-7_5","title":"Generating Contextually Coherent Responses by Learning Structured Vectorized Semantics","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Semantics (computer science); Computer science; Linguistics; Psychology; Natural language processing; Communication; Programming language; Philosophy","score_opus":0.021901773331410506,"score_gpt":0.25503664621272787,"score_spread":0.23313487288131737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3150524533","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028660763,0.000266087,0.96272975,0.00034789884,0.00019215948,0.00014948551,0.00041933323,0.004984164,0.002250404],"genre_scores_gemma":[0.42515537,0.0003995754,0.563366,0.00040821006,0.00030790956,0.00042606622,0.002816197,0.00088962406,0.006231062],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981622,0.0007846279,0.00007333061,0.000553149,0.00027576467,0.00015091084],"domain_scores_gemma":[0.9968677,0.0021865314,0.00013516562,0.00027837505,0.00042412488,0.00010810085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014775512,0.001386742,0.001013668,0.0009061029,0.00041536705,0.0012290622,0.0014822051,0.0015895973,0.00767057],"category_scores_gemma":[0.0074319597,0.00052330684,0.001132557,0.00092941156,0.0005912419,0.0025663748,0.0021535852,0.002070356,0.003702541],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017837523,0.0005249321,0.0014964988,0.0005567012,0.00019016223,0.00033346104,0.0009003653,0.034445047,0.075619124,0.024810148,0.018322136,0.84101766],"study_design_scores_gemma":[0.00014264298,0.00038702597,0.0007758105,0.00006706979,0.000111072135,0.00020152135,0.0005779835,0.89745384,0.025084984,0.06852352,0.00660967,0.000064834836],"about_ca_topic_score_codex":0.0010039614,"about_ca_topic_score_gemma":0.001831402,"teacher_disagreement_score":0.00767057,"about_ca_system_score_codex":0.00043616784,"about_ca_system_score_gemma":0.0008793151,"threshold_uncertainty_score":0.025660634},"labels":[],"label_agreement":null},{"id":"W3151929433","doi":"10.1162/tacl_a_00360","title":"KEPLER: A Unified Model for Knowledge Embedding and Pre-trained Language Representation","year":2021,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":602,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; HEC Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Kepler; Embedding; Benchmark (surveying); Language model; Representation (politics); Natural language processing; Construct (python library); ENCODE; Artificial intelligence; Programming language","score_opus":0.03147363115435537,"score_gpt":0.3240863137469095,"score_spread":0.29261268259255413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3151929433","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017343387,0.0010194231,0.9703542,0.0007331039,0.00014101183,0.0001401233,0.0016710949,0.0067641637,0.001833591],"genre_scores_gemma":[0.44338354,0.0014137506,0.52432823,0.0010410087,0.0002326722,0.00087198964,0.016266946,0.0009986747,0.011463219],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989231,0.00031487033,0.000081702354,0.00043364716,0.00014295544,0.000103744234],"domain_scores_gemma":[0.9974995,0.0013124248,0.00016704587,0.00054583885,0.0003765058,0.00009871783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016194383,0.0016385411,0.001056161,0.0019528219,0.0005481688,0.0018078871,0.0034889902,0.0019793813,0.0038445247],"category_scores_gemma":[0.008467794,0.00085093087,0.0015986672,0.001977759,0.00081028196,0.006360384,0.0029092887,0.0036855463,0.0026577788],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028600797,0.00030525678,0.0021270397,0.00033517537,0.00023734602,0.00026025754,0.0002888322,0.4736743,0.0045394814,0.020317582,0.023665605,0.47396296],"study_design_scores_gemma":[0.000011916734,0.000024553596,0.000119253804,0.000020776573,0.000021780057,0.000032804117,0.000019724419,0.9861498,0.0011783793,0.010473519,0.001935419,0.000012134307],"about_ca_topic_score_codex":0.0070424457,"about_ca_topic_score_gemma":0.010327704,"teacher_disagreement_score":0.0070424457,"about_ca_system_score_codex":0.0012970454,"about_ca_system_score_gemma":0.0016518781,"threshold_uncertainty_score":0.014002919},"labels":[],"label_agreement":null},{"id":"W3152584608","doi":"10.1145/3404835.3462782","title":"Chatty Goose: A Python Framework for Conversational Search","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Python (programming language); Computer science; Goose; Modular design; Programming language; Scratch; World Wide Web","score_opus":0.058706301301711984,"score_gpt":0.31698452297855834,"score_spread":0.25827822167684633,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152584608","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006431519,0.00045581424,0.52430373,0.0007901705,0.00034052663,0.0006065211,0.02311777,0.42682743,0.017126514],"genre_scores_gemma":[0.16809602,0.00068420474,0.63952476,0.001389236,0.0003287738,0.0034221658,0.05572198,0.09917105,0.031661816],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99860793,0.00043811457,0.00011801076,0.0002627819,0.0003663032,0.00020697988],"domain_scores_gemma":[0.9982558,0.00070916355,0.000115997515,0.00036170476,0.0003074996,0.00024981762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00190261,0.0012674556,0.0009918932,0.0014375835,0.0012025296,0.001698255,0.0029081437,0.0010518022,0.036180876],"category_scores_gemma":[0.009164257,0.0009392864,0.0015478161,0.0012024192,0.0008231986,0.0038641286,0.0032917184,0.00288968,0.022157414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011678902,0.00038642707,0.0037760602,0.0022414187,0.00031831278,0.0005047843,0.001556754,0.02540448,0.014240956,0.065540105,0.6562702,0.22859259],"study_design_scores_gemma":[0.00031125342,0.00014771463,0.0033563026,0.00021949112,0.00008255641,0.0005391261,0.0002624483,0.46275833,0.013391249,0.116988085,0.40161118,0.00033230087],"about_ca_topic_score_codex":0.011129636,"about_ca_topic_score_gemma":0.02155912,"teacher_disagreement_score":0.036180876,"about_ca_system_score_codex":0.0010819214,"about_ca_system_score_gemma":0.0033098622,"threshold_uncertainty_score":0.121037126},"labels":[],"label_agreement":null},{"id":"W3152698349","doi":"10.18653/v1/2021.emnlp-main.230","title":"Masked Language Modeling and the Distributional Hypothesis: Order Word Matters Pre-training for Little","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":182,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Word (group theory); Natural language processing; Artificial intelligence; Word order; Language model; Downstream (manufacturing); Order (exchange); Parametric statistics; Linguistics; Mathematics","score_opus":0.07636691541448742,"score_gpt":0.38622034162426394,"score_spread":0.3098534262097765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152698349","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32950336,0.0012406091,0.6500712,0.0029224993,0.00050738227,0.00017802021,0.0010625193,0.008094657,0.0064197965],"genre_scores_gemma":[0.83006823,0.00027660967,0.15962629,0.0013476348,0.000116368145,0.00022936844,0.003277213,0.0007923447,0.0042659417],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984982,0.00068248133,0.00008020862,0.00051439414,0.00011573195,0.00010901767],"domain_scores_gemma":[0.9919951,0.005831776,0.00019366755,0.0013958872,0.00038590064,0.00019771363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003336436,0.0013912388,0.00075726624,0.00047893162,0.000642718,0.0013469394,0.0013256845,0.0013878603,0.0043026595],"category_scores_gemma":[0.017700795,0.0006649107,0.00082722306,0.00056051416,0.0009945612,0.0046881726,0.0015085139,0.0043040034,0.0024966174],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019108157,0.000808427,0.019880205,0.00063665165,0.00035176516,0.0007194346,0.0013434574,0.17906867,0.09119251,0.020549098,0.02816353,0.6553755],"study_design_scores_gemma":[0.000070769725,0.00032681177,0.0032865894,0.000067104855,0.00006628449,0.00022391166,0.00026301018,0.9231259,0.03162676,0.03601005,0.0048786453,0.000054198354],"about_ca_topic_score_codex":0.0039733904,"about_ca_topic_score_gemma":0.008194104,"teacher_disagreement_score":0.0043026595,"about_ca_system_score_codex":0.00062617305,"about_ca_system_score_gemma":0.0016197698,"threshold_uncertainty_score":0.017644942},"labels":[],"label_agreement":null},{"id":"W3152732297","doi":"10.2200/s01078ed2v01y202002hlt049","title":"Semantic Relations Between Nominals, Second Edition","year":2021,"lang":"en","type":"article","venue":"Synthesis lectures on human language technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; University of Toronto","funders":"","keywords":"Curiosity; Mars Exploration Program; Automatic summarization; Computer science; Class (philosophy); Statement (logic); Recall; Semantics (computer science); Linguistics; Question answering; Natural language processing; Artificial intelligence; Psychology; Astrobiology; Programming language; Social psychology; Philosophy","score_opus":0.027629427682573822,"score_gpt":0.278961434224012,"score_spread":0.25133200654143817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3152732297","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0055676824,0.31978512,0.36152467,0.01827591,0.03442307,0.00009606148,0.008686799,0.0039175292,0.24772325],"genre_scores_gemma":[0.11693522,0.14446104,0.25338733,0.0043705716,0.028567307,0.00037727848,0.030796632,0.0032415725,0.41786304],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9993401,0.000104217215,0.00008025986,0.00020636524,0.00023720415,0.000031801585],"domain_scores_gemma":[0.9989796,0.00050000683,0.000041742263,0.00011459037,0.00032025925,0.000043694068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011464418,0.0011172284,0.0010411647,0.0034032576,0.0008547955,0.0052064555,0.0012335412,0.0012287104,0.023139006],"category_scores_gemma":[0.0026378518,0.0009853006,0.0011430188,0.0033904808,0.001917323,0.008201815,0.0010653203,0.0025299485,0.0087835705],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008271036,0.00006596082,0.00034510932,0.0010322914,0.000062318664,0.00013426907,0.0009452989,0.0017167684,0.0019212222,0.26237485,0.47175878,0.25956044],"study_design_scores_gemma":[0.00001643724,0.000017823799,0.0011085063,0.00041828846,0.00003542519,0.00036348257,0.0002573648,0.002514749,0.00060216855,0.18937151,0.80526984,0.000024436167],"about_ca_topic_score_codex":0.00658307,"about_ca_topic_score_gemma":0.0076713064,"teacher_disagreement_score":0.023139006,"about_ca_system_score_codex":0.0024433355,"about_ca_system_score_gemma":0.0014634654,"threshold_uncertainty_score":0.07740766},"labels":[],"label_agreement":null},{"id":"W3153046263","doi":"10.18653/v1/2021.emnlp-main.168","title":"Neural Path Hunter: Reducing Hallucination in Dialogue Systems via Path Grounding","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; University of Alberta","funders":"Alberta Machine Intelligence Institute; Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Path (computing); Security token; Artificial neural network; Focus (optics); Artificial intelligence; Suite; Graph; Deep neural networks; Machine learning; Theoretical computer science; History","score_opus":0.055657809110537304,"score_gpt":0.3775607902916954,"score_spread":0.3219029811811581,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153046263","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29435977,0.003279842,0.6520181,0.0017573764,0.0003331092,0.00063393946,0.0052728457,0.036036152,0.0063089496],"genre_scores_gemma":[0.69971323,0.00041406034,0.2824313,0.00051145384,0.00009473303,0.00033997593,0.0113562895,0.0006843264,0.0044545345],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99811256,0.0007978775,0.00009269069,0.00057433615,0.00029731196,0.00012518368],"domain_scores_gemma":[0.99218106,0.0052176095,0.00037849028,0.0014478329,0.0005767373,0.00019834601],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030506905,0.001702592,0.00075171527,0.0014142495,0.0007260411,0.0012895168,0.0025685166,0.0018445388,0.0030455827],"category_scores_gemma":[0.014086668,0.00042429086,0.0008198223,0.0007865705,0.0013593731,0.0043690456,0.0041663293,0.0022975162,0.0012589439],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020520363,0.00073881407,0.014669017,0.0017425941,0.0003411685,0.0009658912,0.00316566,0.13105634,0.027203433,0.01077154,0.044622976,0.76267064],"study_design_scores_gemma":[0.00021873467,0.00051346334,0.002225341,0.00007450075,0.000119532066,0.00037526392,0.00094937923,0.9294496,0.019073192,0.03334948,0.013586941,0.00006450953],"about_ca_topic_score_codex":0.0039988896,"about_ca_topic_score_gemma":0.008544529,"teacher_disagreement_score":0.0039988896,"about_ca_system_score_codex":0.00079525117,"about_ca_system_score_gemma":0.001119439,"threshold_uncertainty_score":0.016133785},"labels":[],"label_agreement":null},{"id":"W3153390169","doi":"10.1145/3404835.3463049","title":"Improving Transformer-Kernel Ranking Model Using Conformer and Query Term Independence","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Computer science; Inference; Transformer; Artificial intelligence; Machine learning; Data mining; Engineering","score_opus":0.030819142681811164,"score_gpt":0.25686253328186603,"score_spread":0.22604339060005488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153390169","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16894065,0.0016104502,0.8131459,0.00061585824,0.00017902185,0.00012979344,0.0005338528,0.008588744,0.0062558106],"genre_scores_gemma":[0.88996226,0.00049672805,0.09465974,0.00031045667,0.00010722141,0.00008262009,0.0013351289,0.00038756602,0.012658247],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994199,0.00015435298,0.00003869338,0.00016795781,0.00012479244,0.00009432422],"domain_scores_gemma":[0.9988293,0.00042209073,0.00009357162,0.00025841006,0.00033365248,0.00006291358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001626007,0.0008101171,0.0010654574,0.0008816461,0.0003431079,0.0012083085,0.0016193797,0.000998702,0.0025365695],"category_scores_gemma":[0.0033883709,0.0003222515,0.0008564338,0.00085498486,0.00047909442,0.0032041413,0.0008761207,0.0015650578,0.0020939729],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070046313,0.0004981123,0.003189202,0.00019908803,0.00017897226,0.00014127961,0.0001410888,0.51599616,0.018314956,0.011686503,0.011924186,0.4370301],"study_design_scores_gemma":[0.0000115596,0.000047673962,0.00014429279,0.0000025123504,0.000013899872,0.000021585674,0.000005746277,0.9960454,0.0015774039,0.001741909,0.00037948875,0.0000084182775],"about_ca_topic_score_codex":0.012563232,"about_ca_topic_score_gemma":0.01629481,"teacher_disagreement_score":0.012563232,"about_ca_system_score_codex":0.0010014791,"about_ca_system_score_gemma":0.0013751901,"threshold_uncertainty_score":0.024980187},"labels":[],"label_agreement":null},{"id":"W3153765506","doi":"10.18653/v1/2021.eacl-main.263","title":"DISK-CSV: Distilling Interpretable Semantic Knowledge with a Class Semantic Vector","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; University of Alberta; Alberta Machine Intelligence Institute","keywords":"Computer science; Class (philosophy); Natural language processing; Artificial intelligence; Semantic memory; Support vector machine; Information retrieval; Psychology","score_opus":0.015364534898628405,"score_gpt":0.23927037302026563,"score_spread":0.22390583812163722,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153765506","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02277019,0.0004891489,0.95824754,0.00053716905,0.00015469272,0.00016296156,0.0020135073,0.013676873,0.0019479631],"genre_scores_gemma":[0.2493753,0.0005567324,0.73267144,0.00048063253,0.00019180753,0.00032350875,0.010070965,0.0010585171,0.005271175],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992292,0.0001665913,0.000050211627,0.0002784542,0.0002128493,0.00006265598],"domain_scores_gemma":[0.9978472,0.0010525167,0.00016119622,0.00054156006,0.00032044403,0.00007704369],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012086377,0.0014596129,0.00088003336,0.0025072943,0.00070557836,0.0015826683,0.0030683253,0.0017077032,0.005275265],"category_scores_gemma":[0.00579948,0.00044740798,0.001326062,0.0024282467,0.00079437724,0.004424379,0.0026783135,0.0029194981,0.001861123],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034362852,0.00026699234,0.0035989094,0.00036670858,0.00014077775,0.00012175006,0.0005670011,0.046770718,0.0071799867,0.030359965,0.03338489,0.87689865],"study_design_scores_gemma":[0.00005916435,0.00008663153,0.00065581966,0.00005392105,0.000041429335,0.00007541444,0.00017002314,0.9282738,0.0065217353,0.050221775,0.013803868,0.00003636316],"about_ca_topic_score_codex":0.011433394,"about_ca_topic_score_gemma":0.021944124,"teacher_disagreement_score":0.011433394,"about_ca_system_score_codex":0.0011960783,"about_ca_system_score_gemma":0.0022244547,"threshold_uncertainty_score":0.022733688},"labels":[],"label_agreement":null},{"id":"W3153794000","doi":"10.48550/arxiv.2104.05740","title":"A Replication Study of Dense Passage Retriever","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Replication (statistics); Labrador Retriever; Business; Biology; Medicine; Virology; Surgery","score_opus":0.1032666392796188,"score_gpt":0.20511601220934247,"score_spread":0.10184937292972368,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153794000","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.62380683,0.006853176,0.30954143,0.004677685,0.0013495991,0.0030094092,0.0071391277,0.018020801,0.025601955],"genre_scores_gemma":[0.8127333,0.0006589335,0.1643973,0.0015678941,0.00062728353,0.0011557244,0.008662431,0.0008078008,0.009389372],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9885055,0.0068936567,0.0005765496,0.00207136,0.001666627,0.00028634188],"domain_scores_gemma":[0.94771945,0.022956993,0.0010912276,0.021818362,0.0057370085,0.00067699025],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014876824,0.0011006088,0.0017877657,0.0012680806,0.00093512115,0.0015806509,0.0030803797,0.0018362523,0.0058063706],"category_scores_gemma":[0.060316745,0.0005093985,0.001485489,0.0011609105,0.0014104103,0.006236387,0.0022325763,0.0029244495,0.004674888],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0063889544,0.0060748626,0.01497838,0.0031205262,0.0012552124,0.0007939509,0.0026492253,0.07912234,0.052048594,0.027543886,0.077153765,0.72887033],"study_design_scores_gemma":[0.0027214428,0.0090470435,0.016571067,0.00022395917,0.00079905463,0.0013391674,0.0012188944,0.7967954,0.060760338,0.04662157,0.063421175,0.00048094627],"about_ca_topic_score_codex":0.007010292,"about_ca_topic_score_gemma":0.0036172178,"teacher_disagreement_score":0.98512316,"about_ca_system_score_codex":0.001141998,"about_ca_system_score_gemma":0.001277837,"threshold_uncertainty_score":0.07867712},"labels":[],"label_agreement":null},{"id":"W3153874716","doi":"10.48550/arxiv.2104.07058","title":"Predicting Discourse Trees from Transformer-based Neural Summarizers","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Transformer; Natural language processing; Artificial intelligence; Discourse analysis; Style (visual arts); Dependency (UML); Linguistics","score_opus":0.06890917542870136,"score_gpt":0.19568300111540654,"score_spread":0.1267738256867052,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153874716","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47179034,0.002963005,0.5094274,0.0007700186,0.00016979032,0.00017909556,0.0055949437,0.005052013,0.00405342],"genre_scores_gemma":[0.90456957,0.0005406025,0.08342214,0.000058930986,0.000085285916,0.00008641584,0.008663302,0.00012509359,0.0024487092],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995492,0.00018334863,0.0000275535,0.00014288496,0.000055578235,0.000041539326],"domain_scores_gemma":[0.9965863,0.0024554038,0.00022105865,0.00016799413,0.00049476203,0.00007434404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010661874,0.0007314065,0.00043504004,0.0014611781,0.00022626194,0.00070218445,0.0006147303,0.000584838,0.0013225437],"category_scores_gemma":[0.006898165,0.0002274533,0.0004926959,0.000844115,0.00017205892,0.001436475,0.00041766302,0.000915667,0.0009227717],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009014232,0.00024435247,0.014444938,0.00079183414,0.0002609296,0.00024278664,0.0010120341,0.27776203,0.028734092,0.009395488,0.015082322,0.65112776],"study_design_scores_gemma":[0.00002626513,0.00011183229,0.002504626,0.00003380125,0.00006342916,0.00004084202,0.00010821094,0.9784589,0.009237289,0.0074958922,0.0019060924,0.0000127991625],"about_ca_topic_score_codex":0.0022428925,"about_ca_topic_score_gemma":0.0071189036,"teacher_disagreement_score":0.0022428925,"about_ca_system_score_codex":0.0006810273,"about_ca_system_score_gemma":0.00046437173,"threshold_uncertainty_score":0.005638659},"labels":[],"label_agreement":null},{"id":"W3153895709","doi":"10.18653/v1/2021.eacl-srw.21","title":"TMR: Evaluating NER Recall on Tough Mentions","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Recall; Natural language processing; Linguistics; Philosophy","score_opus":0.10726398977631811,"score_gpt":0.35393794426657144,"score_spread":0.2466739544902533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3153895709","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6429636,0.0070865336,0.28201702,0.000836933,0.00067887455,0.00050278037,0.012119648,0.024642821,0.029151846],"genre_scores_gemma":[0.8622268,0.00089524314,0.10655212,0.00023511764,0.00027843475,0.00026935886,0.019336445,0.0015287254,0.008677726],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9927239,0.0019342281,0.0010386183,0.0018698089,0.002044754,0.00038875974],"domain_scores_gemma":[0.9814317,0.009691283,0.0017436086,0.0029974263,0.0038085077,0.00032740197],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011279376,0.0020541043,0.0014590827,0.009425928,0.0008767295,0.002825708,0.0019274886,0.0021664195,0.002215579],"category_scores_gemma":[0.028086001,0.0004462456,0.0011165385,0.0047613126,0.0009486606,0.004378107,0.0018838126,0.0011416657,0.0027019265],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013741312,0.000532524,0.15081854,0.00204392,0.0019250285,0.0006325387,0.0025974184,0.055278633,0.03989586,0.00515013,0.04378991,0.6959613],"study_design_scores_gemma":[0.00024062702,0.0020950185,0.21518213,0.00046774303,0.0015088982,0.003392872,0.0024228322,0.5561141,0.16072877,0.014363688,0.04272251,0.0007609046],"about_ca_topic_score_codex":0.0062095816,"about_ca_topic_score_gemma":0.01049065,"teacher_disagreement_score":0.9887206,"about_ca_system_score_codex":0.00073920825,"about_ca_system_score_gemma":0.0006411612,"threshold_uncertainty_score":0.059651732},"labels":[],"label_agreement":null},{"id":"W3154476213","doi":"10.18653/v1/2021.cmcl-1.9","title":"TorontoCL at CMCL 2021 Shared Task: RoBERTa with Multi-Stage Fine-Tuning for Eye-Tracking Prediction","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Canadian Institute for Advanced Research","keywords":"Computer science; Transformer; Eye tracking; Task (project management); Ranking (information retrieval); Artificial intelligence; Comprehension; Language model; Natural language processing; Tracking (education); Reading comprehension; Task analysis; Machine learning; Reading (process); Programming language; Engineering","score_opus":0.045930620066330856,"score_gpt":0.29172477304761846,"score_spread":0.2457941529812876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154476213","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12600158,0.008958894,0.12738048,0.016850932,0.01722845,0.0040553156,0.36519572,0.29898414,0.03534449],"genre_scores_gemma":[0.16846824,0.000770972,0.14278546,0.0033564363,0.0022539098,0.0040098215,0.59651494,0.016828777,0.06501151],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9908697,0.003295462,0.00036943695,0.0031115422,0.0014537798,0.00090000604],"domain_scores_gemma":[0.9770865,0.005952395,0.00045354982,0.006025375,0.007173074,0.0033091733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012957628,0.007887378,0.0045176647,0.0025289804,0.0030180514,0.005896039,0.006948558,0.008318134,0.028607937],"category_scores_gemma":[0.032362927,0.0018849246,0.003638897,0.0021226571,0.0011962901,0.0047389884,0.007163145,0.006859235,0.047396347],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011231563,0.0007892609,0.0018156334,0.00056177453,0.00031998052,0.00035570262,0.0002541881,0.006978817,0.0061921,0.00071540236,0.89047456,0.0904193],"study_design_scores_gemma":[0.004684852,0.0023787613,0.016258568,0.000490741,0.00066648447,0.0012805981,0.0010155297,0.5083396,0.037270874,0.014406548,0.41249877,0.00070862717],"about_ca_topic_score_codex":0.04522753,"about_ca_topic_score_gemma":0.0747695,"teacher_disagreement_score":0.04522753,"about_ca_system_score_codex":0.0036471442,"about_ca_system_score_gemma":0.0056553134,"threshold_uncertainty_score":0.095703065},"labels":[],"label_agreement":null},{"id":"W3154584133","doi":"10.48550/arxiv.2104.06335","title":"On the Use of Linguistic Features for the Evaluation of Generative Dialogue Systems","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Generalization; Computer science; Relevance (law); Task (project management); Metric (unit); Proposition; Generative grammar; Natural language processing; Artificial intelligence; Measure (data warehouse); Generative model; Machine learning; Linguistics; Mathematics; Data mining","score_opus":0.3171293089957628,"score_gpt":0.24398176982507017,"score_spread":0.07314753917069264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154584133","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55065733,0.0024330942,0.43277526,0.0012593339,0.00014341243,0.0004452712,0.0007976134,0.004186392,0.007302342],"genre_scores_gemma":[0.91184604,0.0001548774,0.086401805,0.00010529157,0.000044721226,0.00014098802,0.0006098518,0.0002269964,0.0004694391],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9744092,0.017456751,0.001038573,0.0028337429,0.0037439028,0.00051786436],"domain_scores_gemma":[0.8626966,0.11923436,0.0048290817,0.0060235625,0.005909124,0.0013072492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026905688,0.0017425242,0.0016420808,0.005100169,0.0011017034,0.0053516384,0.0018513968,0.003325453,0.0012972765],"category_scores_gemma":[0.11695485,0.0005428329,0.0008628314,0.0025538113,0.001684063,0.005001176,0.0029260947,0.0022579483,0.0006380116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022349346,0.00109169,0.051907066,0.0012080318,0.0009059085,0.00028739456,0.0022362736,0.15837371,0.04070558,0.0086365435,0.004284497,0.72812843],"study_design_scores_gemma":[0.00008626113,0.0012876494,0.028378395,0.00019095208,0.000163952,0.0002968213,0.00049270916,0.93478113,0.018048989,0.014390279,0.0017097021,0.00017318569],"about_ca_topic_score_codex":0.0027924771,"about_ca_topic_score_gemma":0.0043357443,"teacher_disagreement_score":0.026905688,"about_ca_system_score_codex":0.0013653381,"about_ca_system_score_gemma":0.0008940277,"threshold_uncertainty_score":0.14229256},"labels":[],"label_agreement":null},{"id":"W3154790171","doi":"10.18653/v1/2021.eacl-main.243","title":"Measuring and Improving Faithfulness of Attention in Neural Machine Translation","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada; Ministère de la Défense Nationale","keywords":"Machine translation; Computer science; Regularization (linguistics); Measure (data warehouse); Artificial intelligence; Translation (biology); Divergence (linguistics); Differentiable function; Deep neural networks; Quality (philosophy); Machine learning; Artificial neural network; Natural language processing; Mathematics; Data mining; Linguistics","score_opus":0.04442249320317228,"score_gpt":0.23021416507049042,"score_spread":0.18579167186731815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154790171","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4020247,0.00072211906,0.59257805,0.00072379655,0.000064970765,0.00007323889,0.000113644666,0.0018009351,0.0018985156],"genre_scores_gemma":[0.95737326,0.000096293945,0.04140254,0.00014019161,0.00003113683,0.000047672915,0.00014187659,0.00018133613,0.00058557483],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99755716,0.0013130008,0.00016728445,0.00049662526,0.00032855227,0.00013732993],"domain_scores_gemma":[0.980824,0.013230734,0.001224412,0.0031195302,0.0012053425,0.00039593488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005696949,0.0010988393,0.0009681899,0.0011250344,0.0006696889,0.001789143,0.0012760091,0.0017081163,0.0012221728],"category_scores_gemma":[0.03782788,0.0006613715,0.00079549866,0.000816849,0.0018560154,0.0035357133,0.0025620917,0.0026543709,0.00026697427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010919324,0.00030044053,0.011984607,0.00027255254,0.00029703006,0.00022823941,0.0010140013,0.74430174,0.027033996,0.018921627,0.0014079759,0.19314578],"study_design_scores_gemma":[0.000025041734,0.00017808525,0.0018948084,0.000019142115,0.000032740474,0.000053393687,0.000049605285,0.9679661,0.0073917415,0.022086635,0.00027818815,0.000024447558],"about_ca_topic_score_codex":0.002511387,"about_ca_topic_score_gemma":0.0026124588,"teacher_disagreement_score":0.005696949,"about_ca_system_score_codex":0.0014378217,"about_ca_system_score_gemma":0.0006851364,"threshold_uncertainty_score":0.030128717},"labels":[],"label_agreement":null},{"id":"W3154793975","doi":"10.1145/3404835.3463120","title":"Vera: Prediction Techniques for Reducing Harmful Misinformation in Consumer Health Search","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Misinformation; Computer science; Credibility; Information retrieval; Relevance (law); Ranking (information retrieval); Context (archaeology); Metric (unit); Data science; Computer security","score_opus":0.04702965344254133,"score_gpt":0.32185475084860915,"score_spread":0.2748250974060678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154793975","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21708472,0.009532049,0.70567256,0.0042296094,0.0010094607,0.0010288209,0.004107203,0.04399586,0.013339681],"genre_scores_gemma":[0.7319327,0.0016224579,0.24195987,0.001013806,0.0007237368,0.00037299076,0.0068285065,0.0011215487,0.014424309],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981846,0.00073091127,0.00011434204,0.00039443397,0.0004283182,0.00014730588],"domain_scores_gemma":[0.99039686,0.006456738,0.00045152436,0.0011152088,0.0013305254,0.00024920187],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037604836,0.0017528023,0.0013163816,0.0036469083,0.00088314805,0.0015050862,0.0017256653,0.0019843641,0.003590916],"category_scores_gemma":[0.015375819,0.0005510516,0.0013766742,0.0018456929,0.00067161984,0.0046105306,0.0013325708,0.002673572,0.0034065845],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012231712,0.0009764932,0.010085103,0.0006137083,0.00026050286,0.00024392916,0.0006569972,0.068146825,0.014023399,0.005922411,0.045657672,0.8521899],"study_design_scores_gemma":[0.00008522337,0.00039028545,0.0015366607,0.000038967308,0.0001226827,0.00019523199,0.00011629973,0.9746629,0.0071956487,0.010400339,0.005212441,0.000043247383],"about_ca_topic_score_codex":0.010148804,"about_ca_topic_score_gemma":0.018084986,"teacher_disagreement_score":0.010148804,"about_ca_system_score_codex":0.0010154204,"about_ca_system_score_gemma":0.0020411946,"threshold_uncertainty_score":0.02017945},"labels":[],"label_agreement":null},{"id":"W3154800935","doi":"10.18653/v1/2021.eacl-main.93","title":"Discourse-Aware Unsupervised Summarization for Long Scientific Documents","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Automatic summarization; Computer science; Exploit; Sentence; Graph; Artificial intelligence; Ranking (information retrieval); Natural language processing; Information retrieval; Topic model; Representation (politics); Machine learning; Theoretical computer science","score_opus":0.027886626085966884,"score_gpt":0.2931997714278539,"score_spread":0.26531314534188705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154800935","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030470049,0.0014344101,0.95823133,0.0004453429,0.00010478882,0.00018718076,0.0013895318,0.0060037877,0.0017337075],"genre_scores_gemma":[0.3546152,0.0012301181,0.62001944,0.00024154999,0.0006243536,0.00062223384,0.013651185,0.0009034607,0.008092533],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990181,0.00033524568,0.000079870755,0.00027432086,0.00021748674,0.00007491645],"domain_scores_gemma":[0.997255,0.0011615154,0.0004902373,0.00025515482,0.000754151,0.000083862586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001124736,0.0013212507,0.0009598964,0.0038373452,0.0005991011,0.0012990148,0.0013295886,0.0008652268,0.0015780479],"category_scores_gemma":[0.0044617793,0.0003790626,0.0008449898,0.0024367198,0.0003431369,0.0020310248,0.00082303624,0.0011813708,0.001736568],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003858433,0.00038765094,0.0027955489,0.0010443788,0.00030379172,0.0002883636,0.0008679551,0.09126769,0.058656257,0.010714003,0.021996042,0.8112924],"study_design_scores_gemma":[0.000051090905,0.00022347919,0.0022986897,0.00005824757,0.00016174836,0.00010401493,0.00020387911,0.94040686,0.025266998,0.019057484,0.012115727,0.000051677278],"about_ca_topic_score_codex":0.0030364476,"about_ca_topic_score_gemma":0.00860074,"teacher_disagreement_score":0.0038373452,"about_ca_system_score_codex":0.00072178873,"about_ca_system_score_gemma":0.001322244,"threshold_uncertainty_score":0.0060375333},"labels":[],"label_agreement":null},{"id":"W3154971029","doi":"10.18653/v1/2021.eacl-main.8","title":"BERxiT: Early Exiting for BERT with Better Fine-Tuning and Extension to Regression","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Waterloo","funders":"Vector Institute; Natural Sciences and Engineering Research Council of Canada; Government of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Inference; Fine-tuning; Extension (predicate logic); Acceleration; Code (set theory); Quality (philosophy); Look-ahead; Regression; Machine learning; Artificial intelligence; Algorithm; Programming language","score_opus":0.0309562390184768,"score_gpt":0.2613005794683832,"score_spread":0.2303443404499064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154971029","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016927667,0.00036462356,0.97241545,0.00026544035,0.000110989145,0.00008662585,0.00010796852,0.008099161,0.001622069],"genre_scores_gemma":[0.39212507,0.0003182464,0.59146154,0.0007175255,0.000264752,0.0003651658,0.0012329861,0.0024450126,0.011069662],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986821,0.0003994937,0.0000748996,0.0003682754,0.00026495528,0.00021021705],"domain_scores_gemma":[0.99657357,0.0017996804,0.0002148091,0.00071489834,0.0004542665,0.00024274737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004264983,0.0017273575,0.001831586,0.00095389015,0.0006713748,0.0016249624,0.0029532206,0.0019286539,0.006320339],"category_scores_gemma":[0.011649348,0.0008799411,0.0011375665,0.0006403887,0.0009001087,0.0032350405,0.0027905551,0.0051013185,0.002577407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010560245,0.0005903226,0.0041859583,0.00030344372,0.00016958266,0.00032735302,0.00035610909,0.40682262,0.017553786,0.021477055,0.015019898,0.5321378],"study_design_scores_gemma":[0.000038944392,0.00007792188,0.00024312848,0.0000149025445,0.00001206732,0.000039690807,0.000012488766,0.9909025,0.0025891804,0.0044171526,0.0016376924,0.000014234979],"about_ca_topic_score_codex":0.0048049437,"about_ca_topic_score_gemma":0.0065446137,"teacher_disagreement_score":0.006320339,"about_ca_system_score_codex":0.00082151266,"about_ca_system_score_gemma":0.0016631419,"threshold_uncertainty_score":0.02255565},"labels":[],"label_agreement":null},{"id":"W3155064524","doi":"10.1007/978-3-030-73103-8_16","title":"Automatic Multiple-Choice and Fill-in-the-Blank Question Generation from Arbitrary Text","year":2021,"lang":"en","type":"book-chapter","venue":"Advances in intelligent systems and computing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Blank; Generator (circuit theory); Computer science; Comprehension; Multiple choice; Artificial intelligence; Natural language processing; Algorithm; Programming language; Linguistics; Reading (process); Power (physics); Engineering","score_opus":0.029862698346860436,"score_gpt":0.2716602425983595,"score_spread":0.24179754425149907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3155064524","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03508252,0.0008481105,0.8950332,0.0008693952,0.0007431411,0.00077846623,0.0043727197,0.05278734,0.009485176],"genre_scores_gemma":[0.19001263,0.0003858101,0.7744262,0.0003386072,0.00037722918,0.0009126592,0.016057665,0.0042509134,0.013238253],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99664634,0.0014815928,0.00023463048,0.0008584931,0.0005385927,0.00024028073],"domain_scores_gemma":[0.98724526,0.009496553,0.00028389075,0.0009171603,0.001756245,0.00030088678],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023801485,0.0023750125,0.0025100245,0.002066124,0.0010847977,0.0023548936,0.0033605334,0.002446485,0.042206213],"category_scores_gemma":[0.012879573,0.0009180855,0.001760138,0.0013440959,0.0008693467,0.004057686,0.0042197537,0.0021246509,0.019784257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012521506,0.0002767476,0.000931202,0.0013155594,0.00008455835,0.000739074,0.001170071,0.006166691,0.05756565,0.016444063,0.07985013,0.8342042],"study_design_scores_gemma":[0.00034485024,0.00035376084,0.0016248358,0.00018137167,0.0001551153,0.0013710853,0.0015209486,0.7623922,0.1077086,0.060915437,0.0632497,0.00018198432],"about_ca_topic_score_codex":0.0011641767,"about_ca_topic_score_gemma":0.0012493184,"teacher_disagreement_score":0.042206213,"about_ca_system_score_codex":0.0008525513,"about_ca_system_score_gemma":0.0010773106,"threshold_uncertainty_score":0.1411938},"labels":[],"label_agreement":null},{"id":"W3155600250","doi":"10.18653/v1/2021.naacl-main.138","title":"Modeling Event Plausibility with Consistent Conceptual Abstraction","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; McGill University","funders":"Canadian Institute for Advanced Research; Compute Canada; Microsoft Research","keywords":"Computer science; Natural language processing; Transformer; Hierarchy; Artificial intelligence; Cognitive psychology; Psychology","score_opus":0.06236140920146336,"score_gpt":0.2777404454129202,"score_spread":0.21537903621145685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3155600250","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034340546,0.00080224377,0.95143753,0.0020547195,0.00010407158,0.000108100656,0.0010791768,0.0012176841,0.008855935],"genre_scores_gemma":[0.7548789,0.0005713403,0.23850633,0.0002989856,0.00017945313,0.00018473616,0.0018586625,0.00028970346,0.003231976],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9947096,0.002439496,0.0003372625,0.0011830578,0.0009817898,0.00034877766],"domain_scores_gemma":[0.97005445,0.022303335,0.0016681644,0.003928647,0.0015142168,0.00053113437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075298175,0.0012143204,0.001284129,0.003535824,0.0013797501,0.005806987,0.003306743,0.0020477714,0.007753827],"category_scores_gemma":[0.05481754,0.0016847798,0.0029166194,0.0026047966,0.0026385428,0.012566534,0.0058298707,0.00407013,0.0011193036],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005142561,0.00009186016,0.004505082,0.0002710583,0.00033065147,0.00064372015,0.0010610282,0.18918754,0.0010375754,0.7439887,0.004831033,0.053537477],"study_design_scores_gemma":[0.000051001192,0.000019678018,0.00028948148,0.000034397977,0.00009279123,0.000104109924,0.00009022843,0.39664587,0.0005227687,0.5988352,0.0032900062,0.000024511912],"about_ca_topic_score_codex":0.008233208,"about_ca_topic_score_gemma":0.009403946,"teacher_disagreement_score":0.008233208,"about_ca_system_score_codex":0.0022854144,"about_ca_system_score_gemma":0.0016983051,"threshold_uncertainty_score":0.039821923},"labels":[],"label_agreement":null},{"id":"W3156463632","doi":"10.1145/3442381.3449977","title":"Typing Errors in Factual Knowledge Graphs: Severity and Possible Ways Out","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Benchmark (surveying); Code (set theory); Artificial intelligence; Typing; Machine learning; Word error rate; Quality (philosophy); Natural language processing; Noisy data; Programming language; Speech recognition","score_opus":0.051265070196393835,"score_gpt":0.2772535504445526,"score_spread":0.22598848024815874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3156463632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40719825,0.0032730638,0.5622968,0.004819508,0.00080299657,0.00036754276,0.005164151,0.010111944,0.0059657623],"genre_scores_gemma":[0.78651726,0.00093182264,0.20072527,0.0006784606,0.0001926086,0.00016326958,0.006319423,0.0021436561,0.002328293],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98021793,0.0074285376,0.0020500435,0.004067284,0.0052740434,0.00096213963],"domain_scores_gemma":[0.7402746,0.19115272,0.016655145,0.031116292,0.018764576,0.0020366532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01562773,0.0011083685,0.001184803,0.0057425667,0.001847397,0.0050044423,0.0025127444,0.0024558883,0.0022085065],"category_scores_gemma":[0.18832901,0.001107031,0.0010119756,0.005404462,0.002699445,0.009648896,0.0032231826,0.0030402653,0.001033659],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012064978,0.00055186625,0.29532582,0.0024876841,0.00067672617,0.0024595307,0.01287384,0.0833008,0.011560639,0.062997036,0.04705109,0.4795085],"study_design_scores_gemma":[0.000100730394,0.00023359645,0.058992,0.0014502736,0.0005184645,0.0042665773,0.007175488,0.6377707,0.030190475,0.21819068,0.040775813,0.00033512723],"about_ca_topic_score_codex":0.006370222,"about_ca_topic_score_gemma":0.010533056,"teacher_disagreement_score":0.01562773,"about_ca_system_score_codex":0.001429678,"about_ca_system_score_gemma":0.0017871547,"threshold_uncertainty_score":0.08264834},"labels":[],"label_agreement":null},{"id":"W3156861296","doi":"10.1145/3404835.3463034","title":"Significant Improvements over the State of the Art? A Case Study of the MS MARCO Document Ranking Leaderboard","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); University of Waterloo","funders":"","keywords":"Ranking (information retrieval); Metric (unit); Context (archaeology); Computer science; State (computer science); sort; Artificial intelligence; Information retrieval; Order (exchange); Machine learning; Algorithm; Engineering; Geography","score_opus":0.02407949558039306,"score_gpt":0.25936462419599887,"score_spread":0.2352851286156058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3156861296","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9294089,0.006197678,0.022643706,0.0066452264,0.00038227925,0.00025007527,0.003199575,0.0028582287,0.028414281],"genre_scores_gemma":[0.9555382,0.0005924931,0.034056615,0.00040927724,0.00032963874,0.000096573815,0.0031685454,0.0006083824,0.0052002487],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9897702,0.0062125046,0.0004048445,0.0010420993,0.0021299347,0.00044049261],"domain_scores_gemma":[0.9579327,0.030186282,0.0017952767,0.004633567,0.0038220643,0.0016300849],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014350804,0.0008303591,0.0010452415,0.0023482386,0.0021582104,0.0037653155,0.0011487171,0.001325884,0.002681391],"category_scores_gemma":[0.042477388,0.00021359776,0.0005753082,0.0034731424,0.0016365227,0.004110443,0.0015518995,0.0020724868,0.0015720333],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0054461025,0.0021480643,0.08959089,0.002260603,0.0007641343,0.002354987,0.01111343,0.054223955,0.019088412,0.033547025,0.14423256,0.6352299],"study_design_scores_gemma":[0.0010865345,0.0075955526,0.24872932,0.0007231378,0.00057874416,0.0028492247,0.02035654,0.36531624,0.047104895,0.06572622,0.23928489,0.0006486862],"about_ca_topic_score_codex":0.0060429005,"about_ca_topic_score_gemma":0.012161466,"teacher_disagreement_score":0.014350804,"about_ca_system_score_codex":0.0014480395,"about_ca_system_score_gemma":0.0007938112,"threshold_uncertainty_score":0.07589519},"labels":[],"label_agreement":null},{"id":"W3157225713","doi":"10.18653/v1/2021.naacl-main.324","title":"Dynabench: Rethinking Benchmarking in NLP","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Office of Naval Research; Defense Advanced Research Projects Agency","keywords":"Benchmarking; Benchmark (surveying); Computer science; Field (mathematics); Artificial intelligence; Data science; Open source; Machine learning; Programming language; Management","score_opus":0.042310958824034024,"score_gpt":0.26710279729644343,"score_spread":0.2247918384724094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157225713","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037155095,0.010026902,0.42230144,0.005558209,0.0045851646,0.0011052996,0.036786716,0.44733194,0.035149287],"genre_scores_gemma":[0.23522691,0.002928016,0.4635672,0.0022556118,0.0007092394,0.0021664943,0.18413414,0.09591017,0.01310223],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97256786,0.014585489,0.0023456523,0.0043053916,0.005199558,0.0009960062],"domain_scores_gemma":[0.96553516,0.020571357,0.00050981855,0.009370897,0.003276282,0.00073664513],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.019682517,0.003269315,0.0021263568,0.0057892115,0.0026117617,0.007945241,0.0070883892,0.0035332278,0.022563852],"category_scores_gemma":[0.08352941,0.0019405512,0.002425446,0.00734752,0.0021374554,0.012720923,0.010784634,0.0047398745,0.016704177],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015511037,0.00059178635,0.006729686,0.0032203358,0.0010119749,0.00048397027,0.0018306848,0.044109263,0.00433333,0.044389952,0.5479149,0.343833],"study_design_scores_gemma":[0.0010666773,0.00032579908,0.003982255,0.00077962555,0.00026993238,0.0002930337,0.0013592036,0.52129304,0.010730103,0.12286147,0.33682156,0.00021719876],"about_ca_topic_score_codex":0.01326843,"about_ca_topic_score_gemma":0.020304035,"teacher_disagreement_score":0.9803175,"about_ca_system_score_codex":0.0022627413,"about_ca_system_score_gemma":0.0036115355,"threshold_uncertainty_score":0.1040923},"labels":[],"label_agreement":null},{"id":"W3157230690","doi":"10.48550/arxiv.2105.00811","title":"CBench: Towards Better Evaluation of Question Answering Over Knowledge Graphs","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Question answering; Benchmark (surveying); Benchmarking; Suite; Information retrieval; Vocabulary; Set (abstract data type); Graph; Syntax; Artificial intelligence; Natural language processing; Task (project management); Theoretical computer science; Programming language; Linguistics","score_opus":0.127953138872214,"score_gpt":0.24176569694518762,"score_spread":0.11381255807297364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157230690","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3857688,0.007632867,0.449221,0.0020852068,0.0011405447,0.0027783872,0.035273485,0.0928095,0.023290187],"genre_scores_gemma":[0.5661734,0.0011850442,0.33685017,0.0008058572,0.00015971712,0.0016514133,0.08536782,0.004291839,0.0035148663],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96368563,0.016297694,0.003614086,0.0037566824,0.011343019,0.0013027918],"domain_scores_gemma":[0.9160568,0.049962748,0.0042177825,0.010276362,0.017211612,0.0022745952],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018426333,0.0030162225,0.0016453745,0.00980896,0.0012065258,0.00415685,0.0045397137,0.0024549284,0.0030248757],"category_scores_gemma":[0.08873054,0.0007156657,0.001613417,0.007663028,0.0015722676,0.005879862,0.0040825075,0.0028958912,0.0014839161],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036008153,0.0037980448,0.041983303,0.01150809,0.001629694,0.00084317045,0.004439589,0.1973083,0.054949135,0.037234575,0.12674661,0.5159586],"study_design_scores_gemma":[0.00039085824,0.0019345806,0.024904625,0.0005272417,0.0002925681,0.0004608628,0.0019962974,0.8301511,0.05233347,0.034662966,0.052114666,0.0002308503],"about_ca_topic_score_codex":0.018391827,"about_ca_topic_score_gemma":0.013875381,"teacher_disagreement_score":0.018426333,"about_ca_system_score_codex":0.0037532395,"about_ca_system_score_gemma":0.0030552754,"threshold_uncertainty_score":0.097448885},"labels":[],"label_agreement":null},{"id":"W3157411108","doi":"10.2196/23898","title":"Health Natural Language Processing: Methodology Development and Applications","year":2021,"lang":"en","type":"editorial","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Harbin Institute of Technology; Natural Science Foundation of Guangdong Province; National Natural Science Foundation of China","keywords":"Scope (computer science); Computer science; Health informatics; Data science; Health care; Natural language; Domain (mathematical analysis); Knowledge management; Natural language processing; Political science","score_opus":0.03536590945122039,"score_gpt":0.377093759670877,"score_spread":0.3417278502196566,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157411108","genre_codex":"commentary","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0007336197,0.11656481,0.27442035,0.32870623,0.26003456,0.0006629977,0.0010633685,0.0022059546,0.015608108],"genre_scores_gemma":[0.013669761,0.14310513,0.30467123,0.05128669,0.46240044,0.0012206136,0.001573151,0.0016445715,0.020428412],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97500676,0.015492091,0.002597497,0.0012444891,0.0054424433,0.00021672703],"domain_scores_gemma":[0.88409424,0.08228846,0.002190692,0.0032635953,0.02669057,0.0014724612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.051522236,0.0010078081,0.0009284808,0.0061489516,0.0011234249,0.0071803993,0.0022603546,0.0035749564,0.004928227],"category_scores_gemma":[0.09750947,0.0008135251,0.0010848669,0.0027452845,0.0037097333,0.004959969,0.0023996616,0.0059499973,0.003247374],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037748356,0.000028838082,0.00019997143,0.0024833314,0.00010699759,0.00014407866,0.0005862543,0.0005464296,0.00047397468,0.023976631,0.69221413,0.27920163],"study_design_scores_gemma":[0.000029360886,0.0000222411,0.000321318,0.0015659146,0.000057725567,0.00019733513,0.00021194441,0.0027664604,0.0005136356,0.04637282,0.9478915,0.000049754817],"about_ca_topic_score_codex":0.0016599811,"about_ca_topic_score_gemma":0.0025629313,"teacher_disagreement_score":0.051522236,"about_ca_system_score_codex":0.002053644,"about_ca_system_score_gemma":0.005476407,"threshold_uncertainty_score":0.27247888},"labels":[],"label_agreement":null},{"id":"W3158244987","doi":"","title":"Evaluating Groundedness in Dialogue Systems: The BEGIN Benchmark","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Benchmark (surveying); Computer science; Metric (unit); Inference; Conversation; Natural language; Artificial intelligence; Natural language processing; Machine learning; Linguistics","score_opus":0.17806079308478368,"score_gpt":0.23581785677706482,"score_spread":0.057757063692281146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158244987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67145973,0.008992088,0.23974136,0.0023974655,0.0009569891,0.003259902,0.030628085,0.017105663,0.025458733],"genre_scores_gemma":[0.829583,0.00059411366,0.11626065,0.0006681909,0.00025485683,0.001848013,0.04600104,0.00079149724,0.0039986256],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98261213,0.01132686,0.0007528585,0.0025815659,0.0021555047,0.00057110883],"domain_scores_gemma":[0.96340984,0.026389549,0.0013882582,0.00493206,0.002565443,0.0013147838],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012380828,0.0024842764,0.0013026809,0.0022891904,0.0012797303,0.001970196,0.0031148246,0.003453407,0.0028963953],"category_scores_gemma":[0.042795926,0.0005201701,0.0012224981,0.0013044626,0.0019512055,0.003496475,0.003971606,0.003286508,0.0017154997],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007614605,0.0062026237,0.03260704,0.0051775803,0.0019080415,0.0009055543,0.0032433954,0.5422639,0.026350945,0.020204302,0.08397353,0.26954857],"study_design_scores_gemma":[0.00090986467,0.0038749708,0.018881975,0.00030794798,0.00023428386,0.0004641002,0.001317086,0.8805325,0.034171727,0.03440341,0.024716731,0.00018532637],"about_ca_topic_score_codex":0.0046786303,"about_ca_topic_score_gemma":0.0075008664,"teacher_disagreement_score":0.012380828,"about_ca_system_score_codex":0.0021154357,"about_ca_system_score_gemma":0.0013550157,"threshold_uncertainty_score":0.065476835},"labels":[],"label_agreement":null},{"id":"W3158589498","doi":"10.18653/v1/2021.nlp4if-1.9","title":"AraStance: A Multi-Country and Multi-Domain Dataset of Arabic Stance Detection for Fact Checking","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; King Abdulaziz City for Science and Technology; Hamad Bin Khalifa University; Compute Canada","keywords":"Disinformation; Computer science; Misinformation; Task (project management); Arabic; Set (abstract data type); Domain (mathematical analysis); Benchmark (surveying); Natural language processing; Artificial intelligence; German; Scale (ratio); Machine learning; Data science; Social media; Computer security; World Wide Web; Linguistics; Mathematics","score_opus":0.05182885206336371,"score_gpt":0.3053056938476754,"score_spread":0.25347684178431173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158589498","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2849635,0.0057082227,0.015560271,0.0025461642,0.00091933564,0.00082087197,0.6481899,0.011782486,0.029509256],"genre_scores_gemma":[0.17057583,0.0009302329,0.031479094,0.00042568397,0.00018557151,0.00040925958,0.7895106,0.0003704137,0.006113264],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99906355,0.00018688435,0.0001254309,0.00021781464,0.00030959386,0.000096843716],"domain_scores_gemma":[0.99674666,0.00109794,0.0003954995,0.0006580384,0.00083494815,0.00026684062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010495476,0.0016790744,0.0005696241,0.0070530376,0.0013149662,0.0014560221,0.0012409177,0.0022100338,0.0066998475],"category_scores_gemma":[0.007391743,0.00026355046,0.00087064446,0.0044461736,0.0004937683,0.0018514647,0.0016384581,0.0014205613,0.0077471016],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00159013,0.0011523137,0.05562009,0.0034984795,0.0004189906,0.0034378497,0.001962267,0.00938047,0.014712412,0.0047076917,0.70698947,0.19652992],"study_design_scores_gemma":[0.00036562548,0.0003228228,0.12135028,0.0008741639,0.0002256576,0.004682888,0.004039553,0.08680776,0.021246124,0.0059294286,0.7539082,0.00024751184],"about_ca_topic_score_codex":0.015085447,"about_ca_topic_score_gemma":0.022913873,"teacher_disagreement_score":0.015085447,"about_ca_system_score_codex":0.00088604377,"about_ca_system_score_gemma":0.0011753304,"threshold_uncertainty_score":0.029995322},"labels":[],"label_agreement":null},{"id":"W3158658358","doi":"10.1109/taslp.2021.3074779","title":"GRTr: Generative-Retrieval Transformers for Data-Efficient Dialogue Domain Adaptation","year":2021,"lang":"en","type":"article","venue":"IEEE/ACM Transactions on Audio Speech and Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"Microsoft Research","keywords":"Computer science; Domain (mathematical analysis); Artificial intelligence; Adaptation (eye); Transformer; Generative grammar; Generative model; Information retrieval; Machine learning; Natural language processing; Mathematics; Engineering","score_opus":0.04292352133080539,"score_gpt":0.29219161997941245,"score_spread":0.24926809864860705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158658358","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013907969,0.00011217623,0.9879488,0.00015416497,0.00004621246,0.000080304024,0.00029660555,0.008085602,0.0018853217],"genre_scores_gemma":[0.24222632,0.00048187783,0.7309318,0.00086904276,0.00024288226,0.00059670873,0.0029373558,0.005046703,0.01666726],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981274,0.00063938316,0.00012380334,0.0005317505,0.00042330427,0.00015439877],"domain_scores_gemma":[0.9972254,0.0013513821,0.00012688618,0.00087850983,0.00029306332,0.0001248996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028146557,0.0013956572,0.0014023393,0.0013372208,0.00066068326,0.0019984134,0.0044440613,0.0019410177,0.02123065],"category_scores_gemma":[0.010982653,0.00091230765,0.0021958181,0.0014043347,0.0018205339,0.0049272776,0.0047851517,0.0035038765,0.012502433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007982295,0.00033239837,0.00083822507,0.0005741155,0.00014045897,0.0006197866,0.0006997831,0.1673203,0.022065887,0.28414914,0.051082425,0.47137922],"study_design_scores_gemma":[0.00006468775,0.00006829378,0.00010096373,0.000021262436,0.000030064013,0.00023205575,0.000042067077,0.8840294,0.007972575,0.09482344,0.012569791,0.00004542835],"about_ca_topic_score_codex":0.0031718083,"about_ca_topic_score_gemma":0.003422556,"teacher_disagreement_score":0.02123065,"about_ca_system_score_codex":0.0013866209,"about_ca_system_score_gemma":0.0013142194,"threshold_uncertainty_score":0.07102358},"labels":[],"label_agreement":null},{"id":"W3158796535","doi":"10.1037/cep0000255","title":"Language experience predicts semantic priming of lexical decision.","year":2021,"lang":"en","type":"article","venue":"Canadian Journal of Experimental Psychology/Revue canadienne de psychologie expérimentale","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg","funders":"","keywords":"Categorization; Priming (agriculture); PsycINFO; Psychology; Natural language processing; Lexical decision task; Semantic memory; Cognition; Semantics (computer science); Cognitive psychology; Computer science; Artificial intelligence; Linguistics; MEDLINE","score_opus":0.042787167974649105,"score_gpt":0.3244604218791198,"score_spread":0.2816732539044707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158796535","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9873886,0.000118343865,0.0034092183,0.00017309187,0.000033545923,0.000053614665,0.00020344758,0.000055465436,0.008564601],"genre_scores_gemma":[0.99678683,0.0000977438,0.0018598192,0.000053079577,0.000025486957,0.00005624099,0.00023342123,0.00004968853,0.00083777343],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99928087,0.00031055612,0.000040412648,0.00022345639,0.000089858484,0.000054991346],"domain_scores_gemma":[0.9773597,0.01866303,0.0021018765,0.0008096597,0.000351677,0.00071407505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002244122,0.00036772448,0.0003635057,0.00061089033,0.00029504742,0.0019725347,0.00034686085,0.0008446854,0.008522295],"category_scores_gemma":[0.025552226,0.00044346802,0.00038957284,0.00057095016,0.0007702392,0.0022071567,0.0010196583,0.0008412336,0.0012694147],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008337588,0.003369548,0.6107431,0.0010930024,0.0005319611,0.0014211276,0.011623357,0.008276247,0.17518178,0.023943264,0.004606118,0.15087284],"study_design_scores_gemma":[0.0003072204,0.0016904571,0.9055959,0.000091355265,0.00033727827,0.00076529395,0.0016459489,0.030848818,0.015877038,0.040698737,0.0020562424,0.00008570248],"about_ca_topic_score_codex":0.0007051754,"about_ca_topic_score_gemma":0.0007048559,"teacher_disagreement_score":0.008522295,"about_ca_system_score_codex":0.00032774487,"about_ca_system_score_gemma":0.00025311968,"threshold_uncertainty_score":0.028509915},"labels":[],"label_agreement":null},{"id":"W3158985534","doi":"","title":"GAIA at SM-KBP 2019 - A Multi-media Multi-lingual Knowledge Extraction and Hypothesis Generation System.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Extraction (chemistry); Natural language processing; Artificial intelligence; Chemistry; Chromatography","score_opus":0.028061753785119277,"score_gpt":0.2686272007403758,"score_spread":0.24056544695525656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158985534","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022626484,0.0015139891,0.4542176,0.0022261806,0.0013271834,0.0016542089,0.09656007,0.4013244,0.018549928],"genre_scores_gemma":[0.08294195,0.00043775642,0.72467417,0.00091776415,0.00032939005,0.00240667,0.15949883,0.012836062,0.015957398],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969823,0.0013030402,0.00022232073,0.000791874,0.0005646166,0.00013577043],"domain_scores_gemma":[0.99319,0.0038581556,0.00023259601,0.0011904176,0.0010885936,0.00044016517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044550006,0.0021291573,0.0016291173,0.0050079175,0.0014410431,0.0033644838,0.003225037,0.0027575057,0.03547538],"category_scores_gemma":[0.017869554,0.0011231005,0.0015973093,0.0023411557,0.00074204593,0.0059480686,0.004255306,0.0026011013,0.03401197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018457621,0.0007536137,0.0028575936,0.0021753993,0.00093655253,0.0017150345,0.0011212578,0.005501966,0.033655256,0.009028254,0.51017356,0.4302358],"study_design_scores_gemma":[0.0017467118,0.0008983151,0.010619206,0.0004800582,0.0006998127,0.0018310952,0.0014344238,0.35561064,0.0613065,0.054570463,0.51031977,0.0004829737],"about_ca_topic_score_codex":0.005095944,"about_ca_topic_score_gemma":0.006427469,"teacher_disagreement_score":0.03547538,"about_ca_system_score_codex":0.0008144039,"about_ca_system_score_gemma":0.0023319377,"threshold_uncertainty_score":0.11867696},"labels":[],"label_agreement":null},{"id":"W3159391311","doi":"","title":"IBM Submission for TAC KBP: EDL 2019 Cascaded Fine-Grained Named Entity Recognition.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"IBM; Computer science; Operating system; Programming language; Artificial intelligence; Materials science; Nanotechnology","score_opus":0.012692283103480919,"score_gpt":0.24858884337689674,"score_spread":0.23589656027341582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159391311","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008278354,0.0020824792,0.061282698,0.011133816,0.017735865,0.00094668707,0.732406,0.08873936,0.07739475],"genre_scores_gemma":[0.013222707,0.00038855744,0.034332715,0.00074320653,0.00097350724,0.0003696985,0.8784395,0.007300037,0.0642301],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99618196,0.0006305917,0.00040849796,0.0008145519,0.0015351237,0.00042923828],"domain_scores_gemma":[0.9869552,0.0022780988,0.00022693806,0.0022201291,0.0070829475,0.0012367667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003951451,0.002522966,0.0026171545,0.004126472,0.0026760306,0.006304191,0.0038919195,0.0033143894,0.2543046],"category_scores_gemma":[0.020224083,0.0014472047,0.00122447,0.0040650265,0.0008614787,0.007545063,0.0047025527,0.0033341837,0.26530483],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001754632,0.000035458957,0.00010264774,0.00016043101,0.000019567791,0.00008021661,0.000024231695,0.00031864163,0.00079318025,0.00089374796,0.98372275,0.013673718],"study_design_scores_gemma":[0.00047008207,0.00012467615,0.001840379,0.00020317791,0.00006218782,0.00038768142,0.00033698315,0.019538218,0.006233978,0.011153157,0.9595406,0.00010888687],"about_ca_topic_score_codex":0.030718025,"about_ca_topic_score_gemma":0.04044182,"teacher_disagreement_score":0.2543046,"about_ca_system_score_codex":0.0023024762,"about_ca_system_score_gemma":0.0044575264,"threshold_uncertainty_score":0.8507336},"labels":[],"label_agreement":null},{"id":"W3159452987","doi":"","title":"The UTexas system for TAC 2019 SM-KBP Task 3: Hypothesis detection with graph convolutional networks.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Graph; Task (project management); Theoretical computer science; Engineering","score_opus":0.006367015457655705,"score_gpt":0.19567945341235674,"score_spread":0.18931243795470104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159452987","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0639654,0.003247601,0.1775468,0.001958538,0.0022669693,0.0015336396,0.20499703,0.51333946,0.031144643],"genre_scores_gemma":[0.2130752,0.0005000019,0.20593573,0.0011346348,0.00045386257,0.0020302967,0.5408706,0.008255656,0.02774393],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988778,0.00024883042,0.00007963579,0.00043695094,0.00021526717,0.00014147135],"domain_scores_gemma":[0.9984894,0.0004346846,0.00008216105,0.00042759013,0.00043084472,0.00013536472],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016036304,0.002289282,0.0011983676,0.002245862,0.0011267762,0.0018995535,0.002807917,0.0023921167,0.02971678],"category_scores_gemma":[0.0070659895,0.0005892891,0.00092090835,0.0012371306,0.00036049262,0.0030076613,0.0026551334,0.0017295897,0.03085798],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015050729,0.00035459388,0.0027515194,0.0007202593,0.00029907623,0.00038956592,0.00024871042,0.006139844,0.012261475,0.0021833056,0.7688257,0.20432083],"study_design_scores_gemma":[0.0011613874,0.00077629375,0.009476962,0.00036147726,0.00035740327,0.0007779264,0.0006659779,0.60886204,0.048047464,0.023657972,0.30562264,0.00023244767],"about_ca_topic_score_codex":0.025782522,"about_ca_topic_score_gemma":0.0345073,"teacher_disagreement_score":0.02971678,"about_ca_system_score_codex":0.0010486781,"about_ca_system_score_gemma":0.002111597,"threshold_uncertainty_score":0.09941256},"labels":[],"label_agreement":null},{"id":"W3159500995","doi":"","title":"Multi-Task Transfer Learning for Fine-Grained Named Entity Recognition.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Task (project management); Transfer of learning; Transfer (computing); Named-entity recognition; Natural language processing; Artificial intelligence; Parallel computing","score_opus":0.016838589117610703,"score_gpt":0.24814881952429116,"score_spread":0.23131023040668044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159500995","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03829663,0.0041593118,0.9407243,0.0012668591,0.000656278,0.00024486706,0.0017796231,0.0086770095,0.004195111],"genre_scores_gemma":[0.74291074,0.0012374927,0.23349388,0.000661993,0.0006814091,0.0005994096,0.008930402,0.00037091013,0.011113683],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99819547,0.0007036693,0.00012534845,0.0005205262,0.00021319641,0.00024182523],"domain_scores_gemma":[0.99554646,0.0023743545,0.00022225644,0.0010864902,0.00055598276,0.00021434492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037032187,0.0012447004,0.0013146395,0.0023182475,0.0011279581,0.0014587527,0.0032864332,0.0023495168,0.0048674243],"category_scores_gemma":[0.0098670265,0.00049175863,0.0012665186,0.0027580669,0.00077455613,0.0053390632,0.0033022498,0.0030999752,0.0036491216],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058118504,0.00070323,0.0023303851,0.00042639955,0.000342523,0.00026052757,0.00030919714,0.07408594,0.0082528535,0.010970478,0.040937632,0.86079955],"study_design_scores_gemma":[0.000030716703,0.0001181912,0.0013155834,0.000029066046,0.00006589288,0.00008453093,0.00016825453,0.9375982,0.0055494453,0.049303927,0.005702488,0.000033634667],"about_ca_topic_score_codex":0.006255192,"about_ca_topic_score_gemma":0.00710878,"teacher_disagreement_score":0.006255192,"about_ca_system_score_codex":0.0010310222,"about_ca_system_score_gemma":0.0014512284,"threshold_uncertainty_score":0.019584715},"labels":[],"label_agreement":null},{"id":"W3159520976","doi":"","title":"The University of Texas at Dallas HLTRI at TAC 2019.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.0059027073133026105,"score_gpt":0.20060592700129407,"score_spread":0.19470321968799145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159520976","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007265395,0.005590681,0.010644045,0.037117645,0.009161728,0.0002645674,0.033281032,0.009285566,0.8873893],"genre_scores_gemma":[0.010604902,0.001284365,0.0040448983,0.001079281,0.00051096204,0.00006011479,0.0064853085,0.00046279837,0.9754674],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991239,0.00010265776,0.000032421023,0.00041376718,0.00019288278,0.00013449359],"domain_scores_gemma":[0.99855226,0.00019493302,0.00006266945,0.00023009896,0.00045143528,0.00050874526],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011274135,0.0011354289,0.0011997832,0.0018346878,0.0020194657,0.0045054513,0.0010267211,0.0014619373,0.65436596],"category_scores_gemma":[0.0018159186,0.00037251055,0.00054208474,0.0021947161,0.00053916656,0.0024417369,0.0022981218,0.0017649016,0.43053156],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022405104,0.000225717,0.0012437414,0.00012785502,0.00002062932,0.00023096288,0.00019282511,0.000197913,0.001793336,0.013812694,0.8359548,0.14597552],"study_design_scores_gemma":[0.00003812926,0.000049896913,0.0018720875,0.00008409796,0.000013263054,0.00010416025,0.00020893203,0.0006358512,0.00058378896,0.0050663543,0.99132663,0.000016775155],"about_ca_topic_score_codex":0.008600288,"about_ca_topic_score_gemma":0.018940223,"teacher_disagreement_score":0.65436596,"about_ca_system_score_codex":0.0027437443,"about_ca_system_score_gemma":0.002839006,"threshold_uncertainty_score":0.49300498},"labels":[],"label_agreement":null},{"id":"W3159905611","doi":"10.18653/v1/2021.nlp4if-1.7","title":"Extractive and Abstractive Explanations for Fact-Checking and Evaluation of News","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; University of Michigan; John Templeton Foundation; National Science Foundation","keywords":"Misinformation; Computer science; Natural language processing; Artificial intelligence; Graph; Information retrieval; Theoretical computer science","score_opus":0.15112999511747213,"score_gpt":0.37202131786644765,"score_spread":0.22089132274897552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159905611","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03995496,0.0012544333,0.9322692,0.0015060416,0.0001542254,0.0007716641,0.0051872605,0.015578964,0.0033232488],"genre_scores_gemma":[0.21520145,0.0003922042,0.7709428,0.00016795503,0.00021160112,0.00038129452,0.01054292,0.0007586766,0.001401077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9839474,0.008852696,0.0011077876,0.0019203468,0.0038410635,0.0003307366],"domain_scores_gemma":[0.8873005,0.08487236,0.0067240754,0.011023072,0.009094308,0.0009857623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012632582,0.0020635838,0.0010971319,0.010825596,0.0011637332,0.0045076115,0.001976012,0.0018835167,0.005789052],"category_scores_gemma":[0.08581514,0.00075789867,0.0016448159,0.0034378835,0.0017196735,0.007998449,0.0025950924,0.0027087862,0.002304682],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011878067,0.0006112031,0.021040868,0.0027720504,0.0004685448,0.00087243656,0.0055209226,0.02640179,0.019001616,0.059269365,0.023447514,0.839406],"study_design_scores_gemma":[0.00032759813,0.00049814506,0.01112312,0.000560079,0.00049308647,0.0009549124,0.0027596068,0.716718,0.07389371,0.12116569,0.071220055,0.0002860339],"about_ca_topic_score_codex":0.0028000134,"about_ca_topic_score_gemma":0.006027083,"teacher_disagreement_score":0.012632582,"about_ca_system_score_codex":0.001530412,"about_ca_system_score_gemma":0.0027155497,"threshold_uncertainty_score":0.06680828},"labels":[],"label_agreement":null},{"id":"W3159905742","doi":"10.48550/arxiv.2010.11140","title":"A Simple and Efficient Multi-Task Learning Approach for Conditioned Dialogue Generation","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Leverage (statistics); Computer science; Exploit; Transformer; Labeled data; Artificial intelligence; Natural language processing; Multi-task learning; Task (project management); Language model; Machine learning","score_opus":0.14099304356237297,"score_gpt":0.2072054085802734,"score_spread":0.06621236501790043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159905742","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007878359,0.0003207939,0.9843739,0.0002122894,0.00012365927,0.0002525364,0.0003002129,0.005336212,0.0012020959],"genre_scores_gemma":[0.38409826,0.0002526894,0.60093755,0.00068671734,0.0003347966,0.0013460309,0.003333017,0.00092612463,0.008084742],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963421,0.0017369693,0.00015279668,0.0011052279,0.00040818588,0.0002547317],"domain_scores_gemma":[0.9937715,0.003712089,0.00024146697,0.00093477295,0.0009893905,0.00035076425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044449745,0.0024741914,0.0017612544,0.001159904,0.00083742104,0.001212418,0.0037792837,0.002418759,0.008487855],"category_scores_gemma":[0.009963264,0.00087509496,0.001837954,0.0010529452,0.0009505334,0.0031623724,0.0035064332,0.004238855,0.0052551185],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013544292,0.001070126,0.0015674912,0.00055192923,0.00024874578,0.00023899596,0.000520267,0.11080996,0.037587713,0.0073494413,0.017911635,0.8207892],"study_design_scores_gemma":[0.00009473045,0.00019384056,0.00045031347,0.000015897767,0.000044122033,0.00009264287,0.00008059551,0.97670823,0.008911936,0.010521687,0.00284156,0.000044459433],"about_ca_topic_score_codex":0.0023420302,"about_ca_topic_score_gemma":0.004317633,"teacher_disagreement_score":0.008487855,"about_ca_system_score_codex":0.00094402034,"about_ca_system_score_gemma":0.002319815,"threshold_uncertainty_score":0.0283947},"labels":[],"label_agreement":null},{"id":"W3160563524","doi":"10.18653/v1/2021.acl-long.114","title":"Reflective Decoding: Beyond Unidirectional Generation with Off-the-Shelf Language Models","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Army Research Office; Naval Information Warfare Center Pacific; Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; Allen Institute for Artificial Intelligence","keywords":"Decoding methods; Computer science; Contextualization; Context (archaeology); Sentence; Task (project management); Artificial intelligence; Reflection (computer programming); Unsupervised learning; Natural language processing; Machine learning; Algorithm","score_opus":0.047130219844225295,"score_gpt":0.2817480599094048,"score_spread":0.23461784006517952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3160563524","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007375424,0.0005189785,0.9751702,0.0007116295,0.00020366041,0.00006821157,0.0002573272,0.00608835,0.009606165],"genre_scores_gemma":[0.39432883,0.0009718495,0.5785494,0.00087327807,0.0003999968,0.000306905,0.0019682008,0.007040844,0.015560703],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997113,0.0015025344,0.00012510951,0.0005635415,0.00047114806,0.00022455459],"domain_scores_gemma":[0.98773164,0.0066162623,0.00022264008,0.0042258473,0.0009558825,0.0002478094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035932355,0.0012572051,0.0012435972,0.0008811833,0.0010562057,0.0051742783,0.0035710502,0.0023554314,0.011702855],"category_scores_gemma":[0.021630656,0.0011011374,0.001253508,0.0012111772,0.001896615,0.010794014,0.006599637,0.002943452,0.009885002],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070267275,0.00028122473,0.0017086423,0.00043525072,0.0001528679,0.0005566715,0.0018124225,0.077834785,0.009998566,0.40899277,0.029635571,0.46788847],"study_design_scores_gemma":[0.00008099759,0.000053574393,0.00007563968,0.000060163205,0.00005192648,0.00017669024,0.00014613337,0.44113508,0.0071944497,0.53490883,0.016079364,0.000037197202],"about_ca_topic_score_codex":0.0022701996,"about_ca_topic_score_gemma":0.003439725,"teacher_disagreement_score":0.011702855,"about_ca_system_score_codex":0.00074208097,"about_ca_system_score_gemma":0.0017814612,"threshold_uncertainty_score":0.03915},"labels":[],"label_agreement":null},{"id":"W3161757350","doi":"10.18653/v1/2021.acl-long.86","title":"MATE-KD: Masked Adversarial TExt, a Companion to Knowledge Distillation","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; University of British Columbia","funders":"","keywords":"Computer science; Distillation; Benchmark (surveying); Divergence (linguistics); Adversarial system; Generator (circuit theory); Artificial intelligence; Key (lock); Set (abstract data type); Test set; Machine learning; Natural language processing; Ranking (information retrieval); Power (physics); Programming language; Linguistics","score_opus":0.03359519427592566,"score_gpt":0.2843402056826809,"score_spread":0.25074501140675526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161757350","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004352532,0.0012606902,0.94801295,0.0020788051,0.0014234338,0.00018194529,0.0028915282,0.028561883,0.011236286],"genre_scores_gemma":[0.21604593,0.0012584287,0.72084296,0.0029413272,0.0012114304,0.0007817801,0.0145116635,0.006650307,0.035756305],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99697185,0.0011353982,0.00014989257,0.0006867738,0.0008528157,0.0002032963],"domain_scores_gemma":[0.99346054,0.0034406625,0.00015658836,0.0021836879,0.0005236773,0.00023483034],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039646174,0.0018529966,0.0019164189,0.001310252,0.0013183006,0.0031838953,0.005076888,0.004158962,0.027702855],"category_scores_gemma":[0.019123564,0.0012719563,0.0013077342,0.0010522372,0.0015913752,0.007621711,0.008963194,0.0053440114,0.01629261],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011083517,0.00046994476,0.0008106215,0.0007437288,0.00034033906,0.0005047422,0.00031398507,0.18045527,0.005758276,0.15299395,0.29232472,0.36417615],"study_design_scores_gemma":[0.00012217226,0.0000759006,0.00013200777,0.000057668334,0.000023273395,0.00016618961,0.00003532878,0.80363905,0.004825258,0.15745635,0.033418473,0.00004831645],"about_ca_topic_score_codex":0.001863166,"about_ca_topic_score_gemma":0.0036099881,"teacher_disagreement_score":0.027702855,"about_ca_system_score_codex":0.0009416541,"about_ca_system_score_gemma":0.0016152783,"threshold_uncertainty_score":0.09267533},"labels":[],"label_agreement":null},{"id":"W3161838415","doi":"10.1007/978-3-030-75765-6_52","title":"SILVER: Generating Persuasive Chinese Product Pitch","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Product (mathematics); Rank (graph theory); Natural language generation; Artificial intelligence; Deep learning; Hierarchy; Statistic; Generator (circuit theory); Artificial neural network; Natural language processing; Information retrieval; Natural language; Power (physics)","score_opus":0.017426941477831692,"score_gpt":0.24929634065845116,"score_spread":0.23186939918061947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161838415","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17254664,0.0013558319,0.66985583,0.0009146151,0.0018101343,0.0010586379,0.0028242688,0.050981287,0.09865279],"genre_scores_gemma":[0.5305633,0.00050332624,0.38318115,0.0003109691,0.00028235134,0.00069653714,0.0051623913,0.0032615745,0.076038554],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996276,0.00010589116,0.000014414518,0.0000910079,0.00013103438,0.000029938556],"domain_scores_gemma":[0.9990477,0.0006181106,0.000028817603,0.0000948316,0.00016514164,0.00004528477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007162976,0.0011659436,0.00050936063,0.00075406465,0.0004920576,0.0010987635,0.0008318611,0.0006961915,0.03408262],"category_scores_gemma":[0.0037144735,0.0003282117,0.00042078816,0.0005867399,0.0002884257,0.0017186579,0.0014425623,0.0006235859,0.0078715095],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088781456,0.00021793964,0.0014924823,0.0007614966,0.00006416715,0.0005053037,0.0019045307,0.008496239,0.03229503,0.013400808,0.08148815,0.85848606],"study_design_scores_gemma":[0.00058672036,0.0018132915,0.0066305622,0.00031478924,0.00031998317,0.00089898636,0.003103151,0.6621862,0.09721779,0.047984693,0.17877528,0.00016852652],"about_ca_topic_score_codex":0.00094302284,"about_ca_topic_score_gemma":0.0011804181,"teacher_disagreement_score":0.03408262,"about_ca_system_score_codex":0.00029580534,"about_ca_system_score_gemma":0.00035152456,"threshold_uncertainty_score":0.114017725},"labels":[],"label_agreement":null},{"id":"W3162439607","doi":"10.31234/osf.io/9a52q","title":"Content matters: Measures of contextual diversity must consider semantic content","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Optimal distinctiveness theory; Diversity (politics); Word (group theory); Computer science; Natural language processing; Word lists by frequency; Context (archaeology); Artificial intelligence; Variance (accounting); Linguistics; Content (measure theory); Psychology; Mathematics; Social psychology; Sociology; History","score_opus":0.25541509696758036,"score_gpt":0.2713124592966478,"score_spread":0.015897362329067466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3162439607","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40536955,0.0038657174,0.5535704,0.002764612,0.00036705757,0.0005165251,0.0022091472,0.0008459236,0.030491007],"genre_scores_gemma":[0.93317556,0.00047193372,0.06317244,0.0004477687,0.0003223927,0.00031138762,0.0011225225,0.00022194101,0.0007539984],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945746,0.002028448,0.00049439096,0.0015093457,0.0010682815,0.0003249597],"domain_scores_gemma":[0.95946145,0.026972668,0.0029172557,0.0068075387,0.0026166688,0.0012243664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006939542,0.0010547191,0.0014402082,0.004365214,0.001447191,0.0047951946,0.0011084619,0.0019067443,0.0032127867],"category_scores_gemma":[0.06105398,0.00061006297,0.0011629061,0.0047989665,0.0032381187,0.0157422,0.0044948813,0.002755318,0.00083327445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014075151,0.0006342411,0.24782263,0.0016858389,0.0017710914,0.00040645347,0.006857827,0.04520267,0.025271747,0.23320442,0.008341383,0.4273942],"study_design_scores_gemma":[0.000108315966,0.0005640865,0.13499863,0.0003160116,0.00046624982,0.00077498605,0.0028339187,0.098449245,0.008978515,0.7341842,0.0180395,0.00028635105],"about_ca_topic_score_codex":0.0024385287,"about_ca_topic_score_gemma":0.0017546279,"teacher_disagreement_score":0.006939542,"about_ca_system_score_codex":0.0012641859,"about_ca_system_score_gemma":0.0010292699,"threshold_uncertainty_score":0.03670025},"labels":[],"label_agreement":null},{"id":"W3162734203","doi":"10.18653/v1/2021.findings-acl.84","title":"A Survey of Data Augmentation Approaches for NLP","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":129,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Popularity; Artificial intelligence; Data science; Natural language processing; Resource (disambiguation); Labeled data; Machine learning; Information retrieval","score_opus":0.5270259338359118,"score_gpt":0.3715261320235383,"score_spread":0.1554998018123735,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3162734203","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027214002,0.047572818,0.9279507,0.0036426522,0.0006694893,0.0003240469,0.0033647907,0.0060022455,0.0077520013],"genre_scores_gemma":[0.06988317,0.07752369,0.8173912,0.0029095723,0.002223541,0.0018203481,0.018327473,0.0021637464,0.0077572283],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99311453,0.0030430453,0.0006312327,0.0015530967,0.0014940829,0.00016407858],"domain_scores_gemma":[0.9840911,0.011811653,0.0004349195,0.0022079556,0.0012917417,0.00016263955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062728617,0.002342331,0.0023906673,0.005542469,0.0014698432,0.0038348397,0.003476076,0.0024897405,0.007733464],"category_scores_gemma":[0.025557682,0.0011401094,0.002967241,0.007716773,0.0019773422,0.0068401066,0.0045376453,0.0043070107,0.0054734885],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013178919,0.00012617385,0.0009920911,0.0032781814,0.0001991499,0.00018344328,0.0005291923,0.017347587,0.0028041643,0.020031769,0.03505859,0.91931784],"study_design_scores_gemma":[0.00008191313,0.00025449554,0.0026673335,0.0032473197,0.00026395146,0.001171967,0.0011579879,0.29558682,0.012514084,0.18052226,0.5023062,0.00022555207],"about_ca_topic_score_codex":0.0040778825,"about_ca_topic_score_gemma":0.0044657784,"teacher_disagreement_score":0.007733464,"about_ca_system_score_codex":0.00140169,"about_ca_system_score_gemma":0.0025844898,"threshold_uncertainty_score":0.033174515},"labels":[],"label_agreement":null},{"id":"W3163735098","doi":"10.48550/arxiv.2105.02923","title":"Hone as You Read: A Practical Type of Interactive Summarization","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Automatic summarization; Computer science; Heuristics; Task (project management); Reading (process); Process (computing); Variety (cybernetics); Metric (unit); Code (set theory); Human–computer interaction; Relevance (law); Artificial intelligence; Programming language","score_opus":0.11909262730609994,"score_gpt":0.24089388922902263,"score_spread":0.12180126192292269,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3163735098","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056037027,0.0015664098,0.890257,0.0014518623,0.00045992585,0.00069393835,0.0042244713,0.03290395,0.012405392],"genre_scores_gemma":[0.23409893,0.0004383677,0.73937994,0.00045741515,0.00032357627,0.0005362014,0.009393353,0.002950261,0.01242188],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99658114,0.0015575818,0.00021325657,0.00088296935,0.00058827305,0.00017675427],"domain_scores_gemma":[0.98919046,0.006258422,0.0005196017,0.002432859,0.0011569213,0.00044182467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031048318,0.0018213339,0.0013878404,0.0014790201,0.0010870321,0.002246563,0.00226446,0.0020380975,0.012425031],"category_scores_gemma":[0.016159853,0.00046086017,0.0009257257,0.001633544,0.00065237744,0.0038970266,0.002675528,0.001721958,0.005566532],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013032529,0.00065369817,0.0021850215,0.0016944301,0.00033574423,0.00043895314,0.0020354998,0.03636121,0.04745235,0.009839096,0.07690667,0.82079405],"study_design_scores_gemma":[0.000500702,0.0021200713,0.004610218,0.00019444592,0.00032563857,0.0012215432,0.0016805182,0.6507573,0.10935913,0.054694023,0.17419414,0.00034228485],"about_ca_topic_score_codex":0.0010958168,"about_ca_topic_score_gemma":0.0033356578,"teacher_disagreement_score":0.012425031,"about_ca_system_score_codex":0.0005572535,"about_ca_system_score_gemma":0.0006813383,"threshold_uncertainty_score":0.041565895},"labels":[],"label_agreement":null},{"id":"W3164193251","doi":"10.48550/arxiv.2105.11698","title":"Guiding the Growth: Difficulty-Controllable Question Generation through Step-by-Step Rewriting","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Rewriting; Interpretability; Controllability; Computer science; Task (project management); Inference; Artificial intelligence; Question answering; Rule of inference; Programming language; Mathematics","score_opus":0.09930973333867954,"score_gpt":0.2001857070874032,"score_spread":0.10087597374872365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3164193251","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03744123,0.0005094103,0.9534749,0.00065100467,0.00006425688,0.0004567165,0.00055525196,0.004507604,0.0023397189],"genre_scores_gemma":[0.2940667,0.0002206568,0.6987284,0.00038149595,0.00007660143,0.00039051136,0.002417224,0.00095220126,0.0027663005],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.993859,0.0033682997,0.00029650636,0.0014310761,0.00082127494,0.0002237663],"domain_scores_gemma":[0.9720514,0.022321546,0.00073979277,0.0029828039,0.0015420002,0.00036247482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055133696,0.0013084238,0.00091782096,0.0011933027,0.0006066129,0.0016288566,0.0024627235,0.001472993,0.0040477044],"category_scores_gemma":[0.03773435,0.0005907177,0.0017660789,0.00074860913,0.0012682956,0.0033751398,0.0029090967,0.002232078,0.0016184788],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077591516,0.0008934,0.014866217,0.001794962,0.00033841323,0.0016404433,0.006795608,0.13283002,0.0800491,0.0791662,0.02368477,0.65716505],"study_design_scores_gemma":[0.00013209516,0.00018512385,0.0014257954,0.000088772664,0.00016160715,0.00043325638,0.000514279,0.84426343,0.029400975,0.10375744,0.019574676,0.00006257145],"about_ca_topic_score_codex":0.0018782342,"about_ca_topic_score_gemma":0.0033260128,"teacher_disagreement_score":0.0055133696,"about_ca_system_score_codex":0.00083719013,"about_ca_system_score_gemma":0.0010649175,"threshold_uncertainty_score":0.029157817},"labels":[],"label_agreement":null},{"id":"W3164690045","doi":"10.14778/3457390.3457398","title":"CBench","year":2021,"lang":"en","type":"article","venue":"Proceedings of the VLDB Endowment","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Question answering; Benchmark (surveying); Benchmarking; Suite; Set (abstract data type); Vocabulary; Information retrieval; Syntax; Graph; Artificial intelligence; Task (project management); Natural language processing; Theoretical computer science; Programming language","score_opus":0.01596362296791629,"score_gpt":0.2186539506523365,"score_spread":0.20269032768442022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3164690045","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03932189,0.004870806,0.20525552,0.0023218761,0.0024000073,0.0025621438,0.084410466,0.41291228,0.24594496],"genre_scores_gemma":[0.1754389,0.0025394442,0.2789048,0.0029252986,0.0003828545,0.0036621531,0.3646212,0.051757623,0.11976772],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99043906,0.00178865,0.0009540799,0.0015757534,0.004353908,0.0008884288],"domain_scores_gemma":[0.9841962,0.003254206,0.00056178164,0.004424116,0.0068580564,0.0007057588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054514743,0.0024405655,0.0014960072,0.0042387904,0.0017061221,0.004822135,0.005067746,0.0021493717,0.054769482],"category_scores_gemma":[0.021978306,0.0009330835,0.0016604166,0.0043366174,0.0010329083,0.005819295,0.0052332585,0.003004603,0.049807444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016898159,0.0007363506,0.0049563395,0.002487053,0.00015006108,0.00040986936,0.0010150596,0.009941691,0.010334826,0.032323323,0.5929157,0.3430398],"study_design_scores_gemma":[0.00023026543,0.0005296072,0.0042209933,0.00035669105,0.00008261956,0.0004934697,0.000688325,0.041283373,0.016735148,0.025531704,0.9096842,0.00016357783],"about_ca_topic_score_codex":0.010205412,"about_ca_topic_score_gemma":0.009110434,"teacher_disagreement_score":0.054769482,"about_ca_system_score_codex":0.0022638135,"about_ca_system_score_gemma":0.0033872374,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3165206559","doi":"10.2196/23099","title":"Predicting Semantic Similarity Between Clinical Sentence Pairs Using Transformer Models: Evaluation and Representational Analysis","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Natural language processing; Semantic similarity; Computer science; Artificial intelligence; Sentence; Security token; Similarity (geometry); Transformer","score_opus":0.15012426007762372,"score_gpt":0.413109033859521,"score_spread":0.26298477378189733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3165206559","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7833918,0.0011971737,0.19360511,0.0014524896,0.0002799409,0.0007649966,0.0040057315,0.011371723,0.003931001],"genre_scores_gemma":[0.9492649,0.00019241805,0.044332728,0.00018036913,0.00003668432,0.00019827932,0.0049175895,0.00018219303,0.00069478934],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99615234,0.0016496491,0.0003503038,0.0009060887,0.0007382353,0.00020325805],"domain_scores_gemma":[0.9809715,0.01476907,0.00065495423,0.0013342138,0.001734006,0.00053627114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00815073,0.0019303382,0.00096490374,0.002384322,0.0006383419,0.0018246232,0.0025601706,0.0016651185,0.002083288],"category_scores_gemma":[0.026734471,0.0005021605,0.0019351799,0.001298823,0.0008473538,0.0030128818,0.0021640828,0.0022976976,0.0009437779],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038081557,0.0013934218,0.055397354,0.0007995337,0.0012307563,0.0007849937,0.00095445296,0.6033831,0.009752919,0.004049203,0.010354183,0.30809197],"study_design_scores_gemma":[0.00003487004,0.00020531508,0.0011816277,0.000014583205,0.00007250607,0.00008619798,0.000071764276,0.99386877,0.0026788963,0.0014710315,0.00029827363,0.000016055239],"about_ca_topic_score_codex":0.017434886,"about_ca_topic_score_gemma":0.014214886,"teacher_disagreement_score":0.017434886,"about_ca_system_score_codex":0.0033685153,"about_ca_system_score_gemma":0.0022728995,"threshold_uncertainty_score":0.043105662},"labels":[],"label_agreement":null},{"id":"W3165961723","doi":"10.18653/v1/2021.acl-long.302","title":"W-RST: Towards a Weighted RST-style Discourse Framework","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Key (lock); Computer science; Style (visual arts); Artificial intelligence; Binary number; Natural language processing; Downstream (manufacturing); Mathematics; Arithmetic; Engineering","score_opus":0.02788330739038415,"score_gpt":0.2980344731926796,"score_spread":0.2701511658022955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3165961723","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00092229946,0.0003411184,0.9947745,0.00035279483,0.000061753635,0.00009626321,0.00054280186,0.0014720703,0.0014364525],"genre_scores_gemma":[0.03314758,0.00057524466,0.95656127,0.00024137713,0.00019601863,0.00035017592,0.0023915272,0.0008528688,0.0056837886],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99295557,0.0035011931,0.0006355757,0.0012448519,0.0013410472,0.00032177308],"domain_scores_gemma":[0.99080265,0.0039782813,0.00068822235,0.0015451788,0.0025366852,0.00044886442],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007586821,0.0018981745,0.0017012415,0.0059835454,0.002276043,0.0061505344,0.0041996688,0.0024312495,0.009863343],"category_scores_gemma":[0.017385958,0.0014341227,0.0028546697,0.0046913135,0.0015847221,0.011743394,0.0057814084,0.003354196,0.006846575],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005237133,0.0002118111,0.0010426264,0.0010128996,0.00027386667,0.00029979195,0.002031269,0.031341664,0.00818338,0.5416522,0.03461627,0.37881064],"study_design_scores_gemma":[0.00007424883,0.00010033543,0.00027924243,0.0002621857,0.00014064524,0.00014040603,0.0004740941,0.36488363,0.0090605505,0.5106792,0.11381284,0.000092710536],"about_ca_topic_score_codex":0.008632732,"about_ca_topic_score_gemma":0.014637284,"teacher_disagreement_score":0.009863343,"about_ca_system_score_codex":0.001663294,"about_ca_system_score_gemma":0.004093163,"threshold_uncertainty_score":0.040123403},"labels":[],"label_agreement":null},{"id":"W3166208174","doi":"10.5715/jnlp.28.350","title":"The Effectiveness of Data Augmentation by Removing Unimportant sentence","year":2021,"lang":"en","type":"article","venue":"Journal of Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Tokyo Metropolitan University; Institute for Catastrophic Loss Reduction","keywords":"Pointer (user interface); Computer science; Generator (circuit theory); Sentence; Natural language processing; Programming language; Artificial intelligence; Physics","score_opus":0.017118238281793714,"score_gpt":0.3082740527003838,"score_spread":0.2911558144185901,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3166208174","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61033845,0.018444834,0.24226132,0.007138699,0.0073184846,0.0010808131,0.020915156,0.06416421,0.028338024],"genre_scores_gemma":[0.7193884,0.0020410055,0.234726,0.0017393441,0.000733923,0.0007416613,0.0275616,0.0012651493,0.01180285],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9964521,0.0014895998,0.00033672486,0.00094110327,0.0005348835,0.0002455381],"domain_scores_gemma":[0.98402184,0.009916619,0.00052281085,0.0031062877,0.0018480953,0.0005843085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049660583,0.003302474,0.0011553278,0.001663545,0.00091088016,0.0018030986,0.0017187995,0.0017227541,0.0051756063],"category_scores_gemma":[0.020553844,0.0007330351,0.001589163,0.00125575,0.0012437825,0.0044386154,0.0021082438,0.0029381951,0.004023359],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00455403,0.00197714,0.0099315895,0.0013457624,0.00042696495,0.0004508857,0.00044274016,0.04812123,0.03591022,0.0026899185,0.06426332,0.82988626],"study_design_scores_gemma":[0.0008208987,0.0024709792,0.010508889,0.0003309722,0.00083370425,0.00074526615,0.0007974905,0.78720385,0.1472121,0.011766784,0.037073076,0.00023603896],"about_ca_topic_score_codex":0.0074228263,"about_ca_topic_score_gemma":0.007638136,"teacher_disagreement_score":0.0074228263,"about_ca_system_score_codex":0.0010100367,"about_ca_system_score_gemma":0.003402674,"threshold_uncertainty_score":0.026263356},"labels":[],"label_agreement":null},{"id":"W3166631396","doi":"10.1613/jair.1.13083","title":"Crossing the Conversational Chasm: A Primer on Natural Language Processing for Multilingual Task-Oriented Dialogue Systems","year":2022,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Task (project management); Pipeline (software); Natural language processing; Focus (optics); Machine translation; Artificial intelligence; Modular design; Conversation; Point (geometry); Field (mathematics); Linguistics; Programming language","score_opus":0.14538255245648596,"score_gpt":0.42541428508142876,"score_spread":0.2800317326249428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3166631396","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010303635,0.12335644,0.83736426,0.012583675,0.0017451718,0.0002745742,0.00018661993,0.0014149994,0.022043804],"genre_scores_gemma":[0.04158784,0.12780799,0.7991263,0.0077922805,0.0057705706,0.0016817976,0.0008354019,0.0014697626,0.0139281],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949768,0.0023065924,0.0005306732,0.00090776524,0.0010645633,0.00021356545],"domain_scores_gemma":[0.99333006,0.0048831785,0.00023451082,0.00060742383,0.0007249819,0.00021992406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008202052,0.0015189709,0.0015912373,0.005153244,0.0013648416,0.00888267,0.004828205,0.005335535,0.0068444167],"category_scores_gemma":[0.009718705,0.0020170559,0.001681662,0.003801266,0.007246613,0.021604666,0.0061055054,0.008615621,0.0046448056],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012797925,0.00018503795,0.0004001303,0.0033385365,0.00008577416,0.00044834908,0.0053986716,0.0046779294,0.0035695634,0.55025667,0.031970397,0.39954093],"study_design_scores_gemma":[0.000019209476,0.00010648696,0.0002825656,0.0022078988,0.0000272516,0.0005361754,0.0009357925,0.018197995,0.0019941386,0.3680974,0.6074804,0.00011472017],"about_ca_topic_score_codex":0.0019256163,"about_ca_topic_score_gemma":0.0013110053,"teacher_disagreement_score":0.00888267,"about_ca_system_score_codex":0.00273982,"about_ca_system_score_gemma":0.0019121926,"threshold_uncertainty_score":0.0433771},"labels":[],"label_agreement":null},{"id":"W3167525829","doi":"10.18653/v1/2021.naacl-main.88","title":"DReCa: A General Task Augmentation Strategy for Few-Shot Natural Language Inference","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Toyota Research Institute; Canadian Institute for Advanced Research","keywords":"Computer science; Inference; Task (project management); Natural language; Computational linguistics; Natural language processing; Artificial intelligence; Shot (pellet); Linguistics; Natural (archaeology); Engineering; Philosophy; History","score_opus":0.049477638520204696,"score_gpt":0.33321349241413817,"score_spread":0.28373585389393347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3167525829","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027698008,0.0007425443,0.950349,0.00029844692,0.00037098562,0.00028162604,0.0019415228,0.040943675,0.0023023328],"genre_scores_gemma":[0.078306206,0.0005506086,0.88901716,0.0008746595,0.0005363679,0.0011818506,0.014520397,0.005212181,0.009800526],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964993,0.001151463,0.00017335088,0.001305257,0.00058972277,0.00028094163],"domain_scores_gemma":[0.99325657,0.0032718813,0.00012273902,0.0020359901,0.0010329266,0.00027982317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050638625,0.003107732,0.0028738135,0.0038711042,0.0024071466,0.0031235076,0.0075690276,0.0032366354,0.020282008],"category_scores_gemma":[0.015657598,0.0020696889,0.0031830373,0.0027574534,0.00094304304,0.0068893605,0.0055366666,0.006963027,0.012548305],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012052492,0.0008921895,0.0014731787,0.0007654312,0.0007875129,0.0002943198,0.0005341782,0.025435057,0.0142218815,0.018213155,0.13691556,0.7992622],"study_design_scores_gemma":[0.00016764912,0.00013096332,0.00060942216,0.000072362,0.00018885553,0.00020325281,0.00012868455,0.89896643,0.010730756,0.05768962,0.031001635,0.00011049547],"about_ca_topic_score_codex":0.01711721,"about_ca_topic_score_gemma":0.046108913,"teacher_disagreement_score":0.020282008,"about_ca_system_score_codex":0.0011257114,"about_ca_system_score_gemma":0.0034098222,"threshold_uncertainty_score":0.06785005},"labels":[],"label_agreement":null},{"id":"W3167799121","doi":"10.18653/v1/2021.naacl-main.326","title":"Predicting Discourse Trees from Transformer-based Neural Summarizers","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Huawei Technologies","keywords":"Automatic summarization; Computer science; Transformer; Natural language processing; Artificial intelligence; Discourse analysis; Dependency (UML); Style (visual arts); Linguistics","score_opus":0.02294131230886106,"score_gpt":0.2563809040675165,"score_spread":0.23343959175865545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3167799121","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4873764,0.0028024493,0.4929314,0.0007049796,0.00016852566,0.00020048105,0.005385978,0.006137184,0.004292555],"genre_scores_gemma":[0.8979391,0.000489443,0.09067205,0.000056848738,0.00007091978,0.00008579149,0.008382528,0.00011302424,0.0021902248],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99956137,0.00018147902,0.000028373668,0.00013677725,0.000051685034,0.00004027569],"domain_scores_gemma":[0.9970878,0.0020740393,0.00019148564,0.00015362477,0.0004267953,0.0000662764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00097414636,0.00073225144,0.0004256208,0.0012809194,0.0002316566,0.00065257703,0.0006062903,0.000591447,0.0013367999],"category_scores_gemma":[0.006098137,0.00022050562,0.00047807212,0.0007258941,0.00016061074,0.0014065587,0.00040274407,0.0009038214,0.0009039992],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008708157,0.00024224866,0.012991851,0.00074333773,0.00025365612,0.00023710208,0.0009146668,0.26937774,0.03077449,0.0072239614,0.012924464,0.66344565],"study_design_scores_gemma":[0.000027518588,0.00012467871,0.0022977418,0.00003175842,0.0000651546,0.00004114598,0.00009738237,0.9803335,0.010029193,0.005236878,0.0017019773,0.000013124784],"about_ca_topic_score_codex":0.0023547066,"about_ca_topic_score_gemma":0.007940889,"teacher_disagreement_score":0.0023547066,"about_ca_system_score_codex":0.0006514765,"about_ca_system_score_gemma":0.0004856053,"threshold_uncertainty_score":0.005151868},"labels":[],"label_agreement":null},{"id":"W3167967639","doi":"","title":"LEGO: Latent Execution-Guided Reasoning for Multi-Hop Question Answering on Knowledge Graphs","year":2021,"lang":"en","type":"article","venue":"International Conference on Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada); University of Alberta","funders":"","keywords":"Computer science; Question answering; Knowledge graph; Hop (telecommunications); Natural language processing; Artificial intelligence; Computer network","score_opus":0.08323708750218646,"score_gpt":0.35839867169081263,"score_spread":0.2751615841886262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3167967639","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009475557,0.00047856875,0.91813296,0.0006091434,0.00013101645,0.00031829,0.004808764,0.063027285,0.0030183268],"genre_scores_gemma":[0.20330645,0.00036178302,0.77036446,0.0005279301,0.00010337942,0.00040320132,0.01658212,0.0032865193,0.0050641503],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982174,0.0005863425,0.00011016009,0.0004786583,0.0004403316,0.00016704164],"domain_scores_gemma":[0.9959777,0.0026364403,0.0001557459,0.00081732287,0.00026733012,0.00014542112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017812711,0.0019679994,0.0013323729,0.0025516404,0.0011243814,0.0032984554,0.0033121142,0.002236361,0.014659976],"category_scores_gemma":[0.009394085,0.0010130556,0.002844297,0.0015103854,0.0010967541,0.0059547145,0.0050665177,0.0030607414,0.004039806],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016125895,0.00093728624,0.0041897255,0.0017492394,0.0006011173,0.0008021699,0.0014748056,0.12433371,0.015719617,0.091834106,0.10571478,0.6510309],"study_design_scores_gemma":[0.00014594938,0.00007671155,0.0003496112,0.000083844374,0.00010205538,0.00007383711,0.00018691705,0.8627472,0.0042833523,0.11988882,0.012022374,0.000039410486],"about_ca_topic_score_codex":0.010565244,"about_ca_topic_score_gemma":0.023124263,"teacher_disagreement_score":0.014659976,"about_ca_system_score_codex":0.0014021238,"about_ca_system_score_gemma":0.002015936,"threshold_uncertainty_score":0.049042463},"labels":[],"label_agreement":null},{"id":"W3168273855","doi":"10.1007/978-3-030-82196-8_60","title":"Neural Abstractive Unsupervised Summarization of Online News Discussions","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Automatic summarization; Computer science; Popularity; Social media; Artificial intelligence; Task (project management); Natural language processing; Information retrieval; Context (archaeology); World Wide Web; Psychology","score_opus":0.02340142679818373,"score_gpt":0.24098635216071476,"score_spread":0.21758492536253105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3168273855","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16606706,0.007360694,0.7787897,0.0019103633,0.0020335575,0.00045237664,0.009704096,0.014017197,0.019664995],"genre_scores_gemma":[0.6540298,0.0024387036,0.25472078,0.00031711115,0.002836702,0.00044777026,0.040590346,0.0012420799,0.04337681],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99927217,0.00020293979,0.00005742042,0.00022000606,0.00013253614,0.00011486806],"domain_scores_gemma":[0.9982778,0.0008653163,0.00014270622,0.00015405401,0.00047699202,0.00008315093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009274746,0.0012961321,0.0011057894,0.0025091032,0.00069749943,0.0017721883,0.0013388653,0.00086858496,0.0060297],"category_scores_gemma":[0.0031847898,0.00046681802,0.0009149139,0.0020561228,0.00027263325,0.0019493196,0.0014292451,0.001443529,0.0041815215],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012426337,0.00031051564,0.0015969821,0.000477525,0.00022033497,0.0001641116,0.0005040969,0.019636301,0.029363057,0.003491248,0.03919697,0.9037962],"study_design_scores_gemma":[0.00008655905,0.00032368483,0.007415804,0.000097468364,0.00030246988,0.00012097243,0.00058827776,0.9314576,0.024960207,0.012659016,0.021930646,0.000057342975],"about_ca_topic_score_codex":0.002910022,"about_ca_topic_score_gemma":0.0067665586,"teacher_disagreement_score":0.0060297,"about_ca_system_score_codex":0.0005731035,"about_ca_system_score_gemma":0.0007287032,"threshold_uncertainty_score":0.020171344},"labels":[],"label_agreement":null},{"id":"W3168291628","doi":"","title":"Aggregating From Multiple Target-Shifted Sources","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Université Laval","funders":"","keywords":"Computer science; Domain adaptation; Domain (mathematical analysis); Artificial intelligence; Similarity (geometry); Adaptation (eye); Conditional probability distribution; Machine learning; Key (lock); Data mining; Selection (genetic algorithm); Multi-source; Mathematics; Statistics","score_opus":0.020207380589618094,"score_gpt":0.2218976012404262,"score_spread":0.2016902206508081,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3168291628","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019465217,0.0006288583,0.9766333,0.00022558446,0.00011255159,0.00008926639,0.00023772834,0.0012311087,0.0013764639],"genre_scores_gemma":[0.5325093,0.0011282614,0.4555404,0.00046807193,0.00044862542,0.00029362706,0.0028785185,0.0005630969,0.006170106],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99836355,0.00045068114,0.00007481412,0.000601864,0.00041387248,0.00009512087],"domain_scores_gemma":[0.99648345,0.001543945,0.00017950046,0.00085113227,0.00082025677,0.000121738114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027174489,0.0022433044,0.0016792926,0.0025122808,0.00079184567,0.0016408485,0.0018054233,0.0012393999,0.0018393253],"category_scores_gemma":[0.0067737764,0.0006319679,0.0018459441,0.0035590255,0.0006849546,0.003072059,0.003391598,0.0023037614,0.0012635388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062415877,0.0004669689,0.009875192,0.0004931949,0.00066216255,0.00087647577,0.0007329622,0.21146946,0.03951316,0.013600956,0.012999645,0.7086856],"study_design_scores_gemma":[0.00004134908,0.00008824831,0.0023459978,0.00004082555,0.00019050603,0.00033575122,0.0002089703,0.94487214,0.01659642,0.028281476,0.006934429,0.00006390422],"about_ca_topic_score_codex":0.0025847293,"about_ca_topic_score_gemma":0.0042225374,"teacher_disagreement_score":0.0027174489,"about_ca_system_score_codex":0.0006553583,"about_ca_system_score_gemma":0.0010082882,"threshold_uncertainty_score":0.014371395},"labels":[],"label_agreement":null},{"id":"W3169035567","doi":"10.1145/3447548.3467196","title":"Reinforced Iterative Knowledge Distillation for Cross-Lingual Named Entity Recognition","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Named-entity recognition; Leverage (statistics); Artificial intelligence; Natural language processing; Ranking (information retrieval); Entity linking; Component (thermodynamics); Benchmark (surveying); Labeled data; Machine learning; Information retrieval; Knowledge base","score_opus":0.0514276670580911,"score_gpt":0.33175434732961856,"score_spread":0.28032668027152746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3169035567","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02243468,0.0008475517,0.96291167,0.00038082452,0.0001662734,0.000114056675,0.0005451619,0.008311502,0.004288335],"genre_scores_gemma":[0.4681301,0.0004329361,0.5150721,0.0007597559,0.00014796724,0.00028311784,0.004992118,0.0008438681,0.009338114],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978879,0.0005805457,0.00013547413,0.00080391526,0.00040324059,0.00018888623],"domain_scores_gemma":[0.99694175,0.001392911,0.00016064975,0.0008794476,0.0005328938,0.00009232644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026154325,0.0012308529,0.0013147302,0.001565156,0.001162849,0.0013328006,0.0030807662,0.0016246754,0.0054824324],"category_scores_gemma":[0.007524256,0.0006151634,0.0012600506,0.0018030222,0.0012328763,0.0045898664,0.0041562086,0.0030385514,0.0030787678],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028978626,0.0003264778,0.0014369771,0.0002738985,0.00016680302,0.00034416252,0.00034204597,0.18707019,0.011784886,0.019778676,0.01372204,0.7644641],"study_design_scores_gemma":[0.000024447532,0.00005679713,0.00026512085,0.000020918302,0.000026518619,0.00006866884,0.000055002805,0.96701473,0.008539622,0.017193181,0.0067041554,0.000030812287],"about_ca_topic_score_codex":0.009699367,"about_ca_topic_score_gemma":0.014736123,"teacher_disagreement_score":0.009699367,"about_ca_system_score_codex":0.0010140941,"about_ca_system_score_gemma":0.0023524673,"threshold_uncertainty_score":0.019285858},"labels":[],"label_agreement":null},{"id":"W3170012708","doi":"10.18653/v1/2021.naacl-main.392","title":"A Simple and Efficient Multi-Task Learning Approach for Conditioned Dialogue Generation","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Leverage (statistics); Transformer; Exploit; Artificial intelligence; Natural language processing; Labeled data; Multi-task learning; Task (project management); Language model; Text generation; Machine learning","score_opus":0.05191918823681962,"score_gpt":0.2668601196632878,"score_spread":0.2149409314264682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3170012708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008690195,0.00034488147,0.9826134,0.00019017796,0.00012266779,0.0002618009,0.0002845526,0.0062223147,0.0012700381],"genre_scores_gemma":[0.39309403,0.00026011563,0.59225774,0.000654348,0.00029339816,0.001262491,0.003132959,0.0009083448,0.008136626],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99690765,0.0013743016,0.00013287783,0.0009678401,0.00038094487,0.00023637079],"domain_scores_gemma":[0.9946911,0.0030587397,0.00021843093,0.00080009206,0.0009262462,0.0003054644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003963672,0.0024397261,0.0016647816,0.0011164135,0.0007581901,0.0011207942,0.0035916532,0.0022266135,0.008526886],"category_scores_gemma":[0.008665948,0.0008127801,0.0017205414,0.0009783417,0.00086752046,0.003100211,0.0032281366,0.0039763525,0.005063317],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012296397,0.0009881315,0.0013475722,0.0005051271,0.00022365575,0.00022242162,0.00046464484,0.09317053,0.04236869,0.0058718137,0.015507803,0.83809996],"study_design_scores_gemma":[0.00010237865,0.00022397323,0.00048639483,0.00001592369,0.000048074053,0.000101023434,0.00008858774,0.97578865,0.011067027,0.008928621,0.0031018686,0.000047489055],"about_ca_topic_score_codex":0.002297737,"about_ca_topic_score_gemma":0.004304443,"teacher_disagreement_score":0.008526886,"about_ca_system_score_codex":0.0008664779,"about_ca_system_score_gemma":0.0022070494,"threshold_uncertainty_score":0.028525293},"labels":[],"label_agreement":null},{"id":"W3171847983","doi":"10.18653/v1/2021.naacl-main.85","title":"Posterior Differential Regularization with f-divergence for Improving Model Robustness","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Robustness (evolution); Computer science; Regularization (linguistics); Computational linguistics; Library science; Natural language processing; Computational biology; Artificial intelligence; Biology; Genetics; Gene","score_opus":0.019840976451395213,"score_gpt":0.22309296642174148,"score_spread":0.20325198997034627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3171847983","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007432517,0.000680188,0.989139,0.00051046244,0.00008929088,0.000037678234,0.00012252,0.0010658118,0.0009225179],"genre_scores_gemma":[0.40302938,0.0010337578,0.58215225,0.0013851215,0.0005795268,0.00037811123,0.0023004888,0.0022662631,0.006875068],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99568295,0.002389707,0.00021762482,0.0008187694,0.0006451867,0.00024576354],"domain_scores_gemma":[0.9895946,0.0076132547,0.0003572079,0.0012640018,0.00088572636,0.00028531387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010828684,0.0025857044,0.0033293022,0.001997129,0.0015271432,0.0023293323,0.0045330836,0.0040765232,0.003608683],"category_scores_gemma":[0.0293826,0.0014049902,0.0022336897,0.0015230676,0.0022379104,0.0038590922,0.0056871986,0.0057144295,0.0015356898],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070145674,0.00022618966,0.0021415849,0.00035412374,0.0006064971,0.00024355915,0.0003371381,0.69458205,0.0067505254,0.052279945,0.018334836,0.22344217],"study_design_scores_gemma":[0.000017328006,0.00001637586,0.0000675571,0.000009997053,0.000014243003,0.00001573284,0.000006671132,0.98563296,0.00048427912,0.013192948,0.00053450704,0.000007373348],"about_ca_topic_score_codex":0.009223465,"about_ca_topic_score_gemma":0.008202433,"teacher_disagreement_score":0.010828684,"about_ca_system_score_codex":0.0019000462,"about_ca_system_score_gemma":0.00259901,"threshold_uncertainty_score":0.057268262},"labels":[],"label_agreement":null},{"id":"W3171978172","doi":"10.18653/v1/2021.naacl-main.333","title":"Inductive Topic Variational Graph Auto-Encoder for Text Classification","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"China Scholarship Council","keywords":"Computer science; Autoencoder; Computational linguistics; Natural language processing; Artificial intelligence; Encoder; Graph; Linguistics; Theoretical computer science; Philosophy; Deep learning","score_opus":0.06204689173918598,"score_gpt":0.2811516705133185,"score_spread":0.21910477877413254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3171978172","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015655605,0.0027285249,0.9685961,0.00088061095,0.00038035665,0.00010571499,0.002019356,0.0058756634,0.003758028],"genre_scores_gemma":[0.47165477,0.0019879376,0.47892332,0.0011035459,0.0008341636,0.0004936866,0.014402763,0.002108559,0.028491348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992085,0.0003037393,0.00003624937,0.0002120244,0.00013559574,0.000103935614],"domain_scores_gemma":[0.99813277,0.00113945,0.0000690992,0.00030644826,0.000280412,0.0000717536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001189677,0.00082761806,0.0013394189,0.0016479824,0.0006523056,0.0009633578,0.0021693974,0.0013428087,0.005128938],"category_scores_gemma":[0.003487933,0.0005882232,0.0011808572,0.0020849868,0.00065956876,0.0026613798,0.0015629668,0.0023049037,0.0032944854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040710723,0.0003202435,0.001633755,0.00031406948,0.00025321727,0.00016629136,0.0002464565,0.13385536,0.00749327,0.06708543,0.08938178,0.6988431],"study_design_scores_gemma":[0.00002442279,0.000024946536,0.00016507873,0.000018621886,0.0000294581,0.000035658348,0.000022780288,0.956027,0.0017122775,0.037718963,0.004210423,0.000010311411],"about_ca_topic_score_codex":0.009252745,"about_ca_topic_score_gemma":0.027880367,"teacher_disagreement_score":0.009252745,"about_ca_system_score_codex":0.0011351228,"about_ca_system_score_gemma":0.0015014089,"threshold_uncertainty_score":0.018397808},"labels":[],"label_agreement":null},{"id":"W3172136722","doi":"10.18653/v1/2021.acl-long.144","title":"Causal Analysis of Syntactic Agreement Mechanisms in Neural Language Models","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation; National Science Foundation","keywords":"Agreement; Sentence; Computer science; Verb; Natural language processing; Artificial intelligence; Linguistics; Subject (documents); Sentence processing; Syntax","score_opus":0.03695777655512643,"score_gpt":0.2821153186159529,"score_spread":0.24515754206082646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3172136722","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19837315,0.0017521516,0.7807007,0.00529548,0.00018106378,0.000052697942,0.0005354558,0.001048451,0.0120608965],"genre_scores_gemma":[0.966223,0.00056925515,0.02920825,0.00015886156,0.00015072401,0.00008280494,0.00032580164,0.00026389142,0.003017391],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988932,0.00064877793,0.000054004922,0.00017519602,0.00013687236,0.0000919775],"domain_scores_gemma":[0.9760439,0.02079959,0.00087587535,0.000951591,0.00091893086,0.00041003092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040866723,0.00047944413,0.0011545665,0.0015734709,0.00092948275,0.0022815014,0.0019523713,0.0013987209,0.007100613],"category_scores_gemma":[0.027783463,0.0009878451,0.0012195243,0.001080868,0.0014812646,0.0055890502,0.0020600127,0.0024478436,0.000543422],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024130278,0.000084234394,0.0033114161,0.00014321525,0.00018396044,0.00018442135,0.00042751344,0.17963305,0.0018046469,0.7851173,0.0033768564,0.02549218],"study_design_scores_gemma":[0.000021173964,0.00001109165,0.00042599186,0.00001146178,0.000029448598,0.000023336941,0.0000318493,0.60267156,0.00029338448,0.39609158,0.00037358233,0.000015561209],"about_ca_topic_score_codex":0.0040746247,"about_ca_topic_score_gemma":0.0045966906,"teacher_disagreement_score":0.007100613,"about_ca_system_score_codex":0.0014252628,"about_ca_system_score_gemma":0.0008443169,"threshold_uncertainty_score":0.023753941},"labels":[],"label_agreement":null},{"id":"W3172246306","doi":"10.48550/arxiv.2004.02251","title":"Semantics of the Unwritten: The Effect of End of Paragraph and Sequence Tokens on Text Generation with GPT2","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Paragraph; Perplexity; Computer science; Artificial intelligence; Natural language processing; Linguistics; Sequence (biology); Semantics (computer science); Language model; Philosophy; Chemistry","score_opus":0.07345739299633469,"score_gpt":0.183247074308855,"score_spread":0.10978968131252032,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3172246306","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8405178,0.0022887462,0.13457155,0.0011141254,0.00044784686,0.00025614267,0.0012930683,0.01042994,0.009080708],"genre_scores_gemma":[0.9562025,0.00021866366,0.037408784,0.00028520616,0.00006913689,0.0001675953,0.0024544138,0.00068802066,0.0025055332],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985672,0.0007130821,0.000088455985,0.00038783616,0.00014984822,0.000093466544],"domain_scores_gemma":[0.9904116,0.0074904333,0.0003286052,0.00085912074,0.0006445803,0.0002657902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025609403,0.0014252843,0.00058461673,0.00061410776,0.00052901666,0.0010975746,0.00090359745,0.0010973362,0.002230007],"category_scores_gemma":[0.015105027,0.00045514666,0.00046308953,0.0004692403,0.0005418745,0.0019191405,0.0012110589,0.0020205402,0.0010452783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023835613,0.0006008842,0.02607856,0.0008136664,0.0002557237,0.0011459542,0.001041417,0.26799017,0.07797597,0.0024295459,0.014231524,0.605053],"study_design_scores_gemma":[0.00023402782,0.00073589076,0.007222068,0.00006345887,0.00016469885,0.0004240362,0.00024517646,0.93117195,0.052051384,0.0034476959,0.004183032,0.00005661961],"about_ca_topic_score_codex":0.0034473815,"about_ca_topic_score_gemma":0.0054673115,"teacher_disagreement_score":0.0034473815,"about_ca_system_score_codex":0.00058721047,"about_ca_system_score_gemma":0.0008137019,"threshold_uncertainty_score":0.013543725},"labels":[],"label_agreement":null},{"id":"W3172341633","doi":"10.1145/3442442.3451385","title":"GOAT at the FinSim-2 task: Learning Word Representations of Financial Data with Customized Corpus","year":2021,"lang":"en","type":"article","venue":"Companion Proceedings of the Web Conference 2021","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Rogers Communications (Canada)","funders":"","keywords":"Word2vec; Computer science; Task (project management); Word (group theory); Rank (graph theory); Ontology; Natural language processing; Domain (mathematical analysis); Artificial intelligence; Information retrieval; Embedding; Mathematics","score_opus":0.039623654592813064,"score_gpt":0.2623438661760558,"score_spread":0.22272021158324276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3172341633","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2670334,0.002890918,0.59030944,0.0016791782,0.0013844543,0.0013352555,0.027900325,0.09296685,0.014500194],"genre_scores_gemma":[0.41811648,0.0005871132,0.45697296,0.0010304335,0.0002676853,0.000958885,0.103948936,0.0015542575,0.016563244],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986854,0.00035621072,0.00010059355,0.0005223463,0.00020714207,0.0001282298],"domain_scores_gemma":[0.9982584,0.000657299,0.00008918841,0.00062000466,0.00026165805,0.00011346797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019612433,0.0031654509,0.0010539336,0.002768678,0.0009435571,0.0016275676,0.0016738462,0.002318929,0.0100270705],"category_scores_gemma":[0.0044054994,0.0005336885,0.0020935405,0.0024121522,0.00051337335,0.0046060816,0.0030515695,0.0020425888,0.0072969864],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000769698,0.0010171987,0.008836944,0.0007077467,0.0005367665,0.0006111203,0.00041789128,0.015950497,0.020737534,0.0050255414,0.114606485,0.8307826],"study_design_scores_gemma":[0.00029302356,0.0011871377,0.011038475,0.00015888432,0.00027600382,0.0016692724,0.0013287674,0.8267957,0.04657039,0.030856779,0.0796353,0.00019022709],"about_ca_topic_score_codex":0.0038296247,"about_ca_topic_score_gemma":0.008441052,"teacher_disagreement_score":0.0100270705,"about_ca_system_score_codex":0.0008765369,"about_ca_system_score_gemma":0.0015322887,"threshold_uncertainty_score":0.033543944},"labels":[],"label_agreement":null},{"id":"W3172351783","doi":"10.18653/v1/2021.sigdial-1.18","title":"Improving Unsupervised Dialogue Topic Segmentation with Utterance-Pair Coherence Scoring","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Utterance; Coherence (philosophical gambling strategy); Artificial intelligence; Inference; Segmentation; Natural language processing; Exploit; Relevance (law); Speech recognition; Task (project management); Machine learning","score_opus":0.02639506892366116,"score_gpt":0.2366871093001384,"score_spread":0.21029204037647722,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3172351783","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.103267506,0.0016693103,0.883767,0.00032397584,0.00013407356,0.00021023046,0.0005563956,0.0073465058,0.0027249248],"genre_scores_gemma":[0.7587633,0.0003603695,0.22972393,0.00024695048,0.00029114672,0.00042369476,0.00449223,0.00086824835,0.0048300372],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973449,0.0012312029,0.000114589355,0.0008130993,0.00031110883,0.00018509125],"domain_scores_gemma":[0.9955771,0.002733133,0.00025467717,0.00046139,0.00075549737,0.00021819725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025366975,0.001600485,0.0012818731,0.0014690581,0.0007034728,0.0012175002,0.0014532198,0.0011794621,0.002103527],"category_scores_gemma":[0.00891987,0.0005223829,0.00089853676,0.0010756918,0.00056299986,0.0024779767,0.002211963,0.002094263,0.0017602821],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010862104,0.0006909169,0.009200456,0.00051382696,0.00043150916,0.00020044489,0.002092434,0.07519821,0.057615705,0.004956938,0.014409121,0.83360416],"study_design_scores_gemma":[0.000057086803,0.00019064083,0.0036346272,0.0000258901,0.00010097126,0.00007286926,0.0002760044,0.9779288,0.009946511,0.0051223575,0.0026069107,0.000037314607],"about_ca_topic_score_codex":0.004924356,"about_ca_topic_score_gemma":0.00874891,"teacher_disagreement_score":0.004924356,"about_ca_system_score_codex":0.0005982598,"about_ca_system_score_gemma":0.0014259267,"threshold_uncertainty_score":0.013415515},"labels":[],"label_agreement":null},{"id":"W3172401221","doi":"10.1017/pan.2021.15","title":"Multi-Label Prediction for Political Text-as-Data","year":2021,"lang":"en","type":"article","venue":"Political Analysis","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Centre for Social Innovation","funders":"Compute Canada; National Science Foundation","keywords":"Computer science; Machine learning; Artificial intelligence; Code (set theory); Set (abstract data type); Supervised learning; Training set; Source code; Government (linguistics); Data set; Association (psychology); Natural language processing; Artificial neural network; Psychology","score_opus":0.0953883885612755,"score_gpt":0.35060281141808164,"score_spread":0.2552144228568061,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3172401221","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48619235,0.0023071982,0.48628315,0.0055090063,0.0009017647,0.00031570622,0.008015973,0.0058767158,0.004598022],"genre_scores_gemma":[0.90976506,0.00018247019,0.082297616,0.0002096775,0.00040173184,0.0001532105,0.0058710864,0.0001488445,0.00097033434],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948737,0.0030751047,0.00018406127,0.0009607768,0.0006434192,0.00026299513],"domain_scores_gemma":[0.93579805,0.05084828,0.0042364253,0.0040585254,0.004091955,0.00096666324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008189013,0.000898994,0.0008082141,0.0039976877,0.0011733538,0.0018360504,0.0014283261,0.0015025144,0.002349821],"category_scores_gemma":[0.037890643,0.0002976839,0.0006366962,0.0038071391,0.00083386333,0.0033324647,0.0013857464,0.003500791,0.0014758088],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010158474,0.0015153427,0.15514117,0.0005143407,0.00036351182,0.00048894476,0.0009439923,0.4039959,0.0049028606,0.019041872,0.04565397,0.36642224],"study_design_scores_gemma":[0.000008575915,0.000014952552,0.0030877914,0.000017473752,0.000006404688,0.0000144748155,0.00006177874,0.9851732,0.0007455895,0.009823857,0.0010364015,0.000009480222],"about_ca_topic_score_codex":0.005781762,"about_ca_topic_score_gemma":0.007544234,"teacher_disagreement_score":0.008189013,"about_ca_system_score_codex":0.0012251548,"about_ca_system_score_gemma":0.0008982957,"threshold_uncertainty_score":0.04330814},"labels":[],"label_agreement":null},{"id":"W3173373527","doi":"","title":"TREC 2020 Notebook: CAsT Track.","year":2020,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Track (disk drive); Computer science; Operating system","score_opus":0.06551698452288615,"score_gpt":0.26444549199311357,"score_spread":0.19892850747022742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173373527","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020497588,0.013639513,0.009974361,0.0070614917,0.020928705,0.0017334512,0.72682226,0.020553326,0.1972371],"genre_scores_gemma":[0.005310373,0.0051755463,0.010279756,0.0022997712,0.0043809623,0.0011491418,0.6500084,0.004168643,0.31722742],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99712014,0.0005336733,0.0001765314,0.00028451718,0.0015009204,0.00038425848],"domain_scores_gemma":[0.98296183,0.0015955978,0.0007039419,0.0011857081,0.010203678,0.0033493892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066589573,0.004504459,0.00308029,0.009631332,0.0024312714,0.00719639,0.005421328,0.0023559448,0.26791924],"category_scores_gemma":[0.011002514,0.0013136587,0.0014383333,0.010259045,0.00087799306,0.0061250767,0.0025363849,0.0029561159,0.28209016],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030264131,0.000013667625,0.000019958325,0.00013920518,0.000005262087,0.0000034373434,0.0000020904267,0.00008060235,0.00021406796,0.000121716614,0.9926087,0.006760996],"study_design_scores_gemma":[0.00022707015,0.00017962075,0.0024549435,0.00028850685,0.00006573498,0.000073169474,0.000047826634,0.0016364877,0.0023195746,0.0020724721,0.99052715,0.00010734794],"about_ca_topic_score_codex":0.16751295,"about_ca_topic_score_gemma":0.21497685,"teacher_disagreement_score":0.26791924,"about_ca_system_score_codex":0.004675356,"about_ca_system_score_gemma":0.012770094,"threshold_uncertainty_score":0.89627916},"labels":[],"label_agreement":null},{"id":"W3173471803","doi":"10.18653/v1/2021.findings-acl.335","title":"Named Entity Recognition through Deep Representation Learning and Weak Supervision","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Representation (politics); Natural language processing; Artificial intelligence; Deep learning; Speech recognition; Political science","score_opus":0.04472971849717734,"score_gpt":0.29232644409387615,"score_spread":0.24759672559669882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173471803","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020307116,0.00041780542,0.97475797,0.00029510353,0.000041465108,0.00004453897,0.00033306764,0.0029337623,0.00086926616],"genre_scores_gemma":[0.58948404,0.0007698595,0.39626727,0.00042904582,0.0001527989,0.0003050319,0.006015862,0.00029067873,0.006285406],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99831724,0.0006687491,0.00008911724,0.0005310052,0.00025465118,0.00013917028],"domain_scores_gemma":[0.9964903,0.0014749089,0.0003658937,0.0010662178,0.00049137307,0.00011137778],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002544552,0.0013039222,0.0013104728,0.0013867677,0.0005677176,0.0014783618,0.0028355266,0.0014541743,0.0013555185],"category_scores_gemma":[0.0063021155,0.00058506324,0.0012150129,0.0018112073,0.00095971854,0.005741287,0.0024971433,0.0032000123,0.0016904146],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036064954,0.0003997665,0.0041778525,0.00022485989,0.00020864973,0.0001799409,0.00034025792,0.24033499,0.010151079,0.021819394,0.012945934,0.70885664],"study_design_scores_gemma":[0.000007383914,0.000028218083,0.0003005655,0.000013838688,0.000015617054,0.00002209264,0.000021895139,0.9795304,0.0026462718,0.016487688,0.0009159289,0.000010075419],"about_ca_topic_score_codex":0.003474197,"about_ca_topic_score_gemma":0.0060773552,"teacher_disagreement_score":0.003474197,"about_ca_system_score_codex":0.0009808559,"about_ca_system_score_gemma":0.0009926519,"threshold_uncertainty_score":0.01345706},"labels":[],"label_agreement":null},{"id":"W3173618889","doi":"10.18653/v1/2021.acl-long.84","title":"The Art of Abstention: Selective Prediction and Error Regularization for Natural Language Processing","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computational linguistics; Computer science; Regularization (linguistics); Natural language processing; Joint (building); Artificial intelligence; Linguistics; Philosophy; Engineering","score_opus":0.010330237839250538,"score_gpt":0.24780675794415752,"score_spread":0.23747652010490697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173618889","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008311779,0.00596733,0.9744815,0.007843235,0.0004550674,0.00003615129,0.00015032843,0.00073038024,0.0020241532],"genre_scores_gemma":[0.46783096,0.0061101336,0.50395346,0.0048285862,0.004810073,0.0006209255,0.0007576824,0.0012049571,0.009883234],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916475,0.005439586,0.00040383014,0.001240248,0.0009778539,0.00029097195],"domain_scores_gemma":[0.94753087,0.0444747,0.00078736205,0.0051981513,0.0014810356,0.0005278177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015277454,0.0016392384,0.003156415,0.0017315635,0.0019663563,0.003366747,0.0039181383,0.0034182228,0.0028423246],"category_scores_gemma":[0.047987938,0.0017952576,0.0016345544,0.0022353586,0.006366044,0.01198248,0.0062363134,0.011134302,0.0011593822],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010595465,0.00027874194,0.002625919,0.00076990307,0.00057593995,0.00019730261,0.0009600189,0.17150874,0.003361835,0.35529217,0.052203193,0.41116676],"study_design_scores_gemma":[0.000043890635,0.00003544359,0.00014208096,0.000034620625,0.000028667504,0.000024382616,0.000027752158,0.5185311,0.0005960579,0.47819528,0.0023199,0.000020785841],"about_ca_topic_score_codex":0.0050168186,"about_ca_topic_score_gemma":0.006138278,"teacher_disagreement_score":0.015277454,"about_ca_system_score_codex":0.0014182809,"about_ca_system_score_gemma":0.0016232304,"threshold_uncertainty_score":0.080795884},"labels":[],"label_agreement":null},{"id":"W3173761280","doi":"10.1162/neco_a_01410","title":"Simulating and Predicting Dynamical Systems With Spatial Semantic Pointers","year":2021,"lang":"en","type":"article","venue":"Neural Computation","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; University of Waterloo","funders":"","keywords":"Computer science; Dynamical systems theory; Symbol (formal); Artificial intelligence; Theoretical computer science; Artificial neural network; Task (project management); Range (aeronautics); Exploit; Dynamical system (definition); Space (punctuation); Machine learning; Physics","score_opus":0.015112865175800641,"score_gpt":0.24285977625415267,"score_spread":0.22774691107835204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173761280","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5653305,0.00019919737,0.42997956,0.0007575053,0.00007673761,0.000045109307,0.00030283397,0.0004940749,0.0028144429],"genre_scores_gemma":[0.9515551,0.0001204553,0.047086343,0.000043449923,0.000023046126,0.000054743658,0.0002072882,0.000036947178,0.00087260024],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998079,0.000079488964,0.000011601218,0.000045894212,0.000028778428,0.00002651466],"domain_scores_gemma":[0.99848986,0.0010903948,0.00014960828,0.00012156689,0.00007525309,0.00007327154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007838876,0.00040489523,0.0004488931,0.0005888509,0.00027806356,0.00083039945,0.00072328426,0.0009982507,0.0012032521],"category_scores_gemma":[0.0045289584,0.00035325147,0.00061514846,0.0005993907,0.0010760254,0.0021160417,0.00087349745,0.0009954433,0.00012498043],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018888803,0.000014258072,0.0015812855,0.0000147507235,0.000011529923,0.00002204327,0.000043591604,0.9809731,0.00032913644,0.013456438,0.00015493449,0.003380033],"study_design_scores_gemma":[0.0000019759266,0.0000026983566,0.00006382638,9.054167e-7,8.2924385e-7,0.0000016096037,0.0000054844418,0.9951925,0.000062851526,0.004628926,0.000037228212,0.0000012671679],"about_ca_topic_score_codex":0.006638874,"about_ca_topic_score_gemma":0.008857207,"teacher_disagreement_score":0.006638874,"about_ca_system_score_codex":0.00077597896,"about_ca_system_score_gemma":0.0006587546,"threshold_uncertainty_score":0.013200462},"labels":[],"label_agreement":null},{"id":"W3173783447","doi":"10.18653/v1/2021.acl-long.72","title":"DeCLUTR: Deep Contrastive Learning for Unsupervised Textual Representations","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":362,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; Vector Institute; University of Toronto","funders":"National Institutes of Health; Compute Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Natural language processing; Artificial intelligence; Linguistics; Computational linguistics; Volume (thermodynamics); Association (psychology); Joint (building); Deep learning; Cognitive science; Psychology; Philosophy; Epistemology; Engineering","score_opus":0.031081299079669916,"score_gpt":0.28708953364301465,"score_spread":0.25600823456334476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173783447","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014980346,0.0028195502,0.92518044,0.0010941183,0.0005837438,0.00015680003,0.006015741,0.043328144,0.005841149],"genre_scores_gemma":[0.2472118,0.0014275376,0.6886694,0.0012251335,0.00049185415,0.00060467253,0.03292605,0.0044308226,0.02301267],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999358,0.0002180189,0.00003455754,0.00022286197,0.0001113503,0.000055170258],"domain_scores_gemma":[0.99850935,0.00085367035,0.00006700624,0.0003077256,0.00019477152,0.00006750428],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012636568,0.0013246502,0.0009180878,0.0016788337,0.00071885955,0.0018443107,0.0027685566,0.0013313721,0.008955889],"category_scores_gemma":[0.004383939,0.0008025209,0.0011600554,0.0015900544,0.0006686,0.0049016224,0.0024700195,0.0032733341,0.0064846184],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056970556,0.00033286118,0.0012066539,0.0005123017,0.00022081414,0.00018461804,0.0002644453,0.03556187,0.012613568,0.026352186,0.13826172,0.7839192],"study_design_scores_gemma":[0.00012110368,0.00011150617,0.00052586285,0.00007573316,0.000055843044,0.000107632884,0.00009220035,0.88455564,0.01114882,0.073889844,0.029272193,0.00004367623],"about_ca_topic_score_codex":0.0053960946,"about_ca_topic_score_gemma":0.012424376,"teacher_disagreement_score":0.008955889,"about_ca_system_score_codex":0.0010083623,"about_ca_system_score_gemma":0.00093457376,"threshold_uncertainty_score":0.029960454},"labels":[],"label_agreement":null},{"id":"W3173886731","doi":"10.18653/v1/2021.acl-short.119","title":"Demoting the Lead Bias in News Summarization via Alternating Adversarial Learning","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Huawei Technologies","keywords":"Automatic summarization; Adversarial system; Computer science; Association (psychology); Joint (building); Natural language processing; Artificial intelligence; Linguistics; Engineering; Philosophy; Epistemology","score_opus":0.045812041800633564,"score_gpt":0.2636030402067356,"score_spread":0.21779099840610203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173886731","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023203352,0.0023098623,0.9680636,0.0013616041,0.00038227084,0.00009956717,0.00024534212,0.0015091359,0.0028252874],"genre_scores_gemma":[0.7030453,0.001435282,0.27412683,0.0018101288,0.0018913256,0.00046468028,0.0018718669,0.00097618345,0.014378379],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983298,0.0009166413,0.0000783506,0.00034053004,0.00021035725,0.00012431027],"domain_scores_gemma":[0.9911848,0.006906751,0.00043967634,0.0006347348,0.0006217739,0.00021231532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042111524,0.001913252,0.0018979089,0.0008614324,0.0007401874,0.0014111338,0.0021069862,0.0020475062,0.0035107294],"category_scores_gemma":[0.01420693,0.000936866,0.0007528845,0.0008302254,0.0013024911,0.003094357,0.0025393195,0.003221818,0.0016690666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001097661,0.00018285368,0.0012720989,0.0005211724,0.000287553,0.00019606814,0.00050097296,0.57842267,0.00894205,0.027785132,0.031291723,0.34949994],"study_design_scores_gemma":[0.000040754523,0.00006967149,0.00012593411,0.000021351187,0.000030515786,0.000020255287,0.000025160845,0.9815639,0.0017384667,0.015318385,0.0010333172,0.0000123945565],"about_ca_topic_score_codex":0.0022759433,"about_ca_topic_score_gemma":0.0038410602,"teacher_disagreement_score":0.0042111524,"about_ca_system_score_codex":0.00076276297,"about_ca_system_score_gemma":0.0009529083,"threshold_uncertainty_score":0.022270977},"labels":[],"label_agreement":null},{"id":"W3173964589","doi":"10.1609/aaai.v35i14.17550","title":"Dynamic Hybrid Relation Exploration Network for Cross-Domain Context-Dependent Semantic Parsing","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Parsing; Artificial intelligence; Natural language processing; Context (archaeology); Representation (politics); Graph; Focus (optics); Relation (database); Dependency grammar; Domain (mathematical analysis); Natural language understanding; Conversation; Natural language; Theoretical computer science; Data mining","score_opus":0.09084251400592146,"score_gpt":0.3157267080256453,"score_spread":0.22488419401972382,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173964589","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04584973,0.0019712118,0.9379827,0.000991918,0.0000954507,0.00015698411,0.0016859486,0.0063218824,0.004944169],"genre_scores_gemma":[0.62537247,0.0010582658,0.35498697,0.0007318965,0.00011533708,0.0003821751,0.0066586537,0.000656054,0.010038253],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99879116,0.00034355762,0.000054169515,0.0005248034,0.00018862769,0.0000977649],"domain_scores_gemma":[0.9984302,0.00092650845,0.00009536011,0.0002662757,0.00021040929,0.000071263916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017850277,0.00096224324,0.0009573983,0.0026085677,0.0009124237,0.0012062944,0.0023583097,0.0016068339,0.0035939834],"category_scores_gemma":[0.004950858,0.00060049404,0.0012755078,0.0028420882,0.0009266002,0.004509495,0.0023009344,0.001864848,0.0011842765],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006688091,0.00041020542,0.00677051,0.0003612232,0.00024709446,0.00072423136,0.0011662863,0.38028434,0.00908223,0.072718486,0.022553502,0.50501317],"study_design_scores_gemma":[0.00001425756,0.000025446001,0.00043727166,0.000019725312,0.00004015802,0.0000873503,0.000069436675,0.9567845,0.0015499152,0.03633885,0.0046140915,0.000018965886],"about_ca_topic_score_codex":0.01754799,"about_ca_topic_score_gemma":0.028987622,"teacher_disagreement_score":0.01754799,"about_ca_system_score_codex":0.0014847734,"about_ca_system_score_gemma":0.002078714,"threshold_uncertainty_score":0.034891725},"labels":[],"label_agreement":null},{"id":"W3173999950","doi":"10.1609/aaai.v35i16.17654","title":"On Scalar Embedding of Relative Positions in Attention Models","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Beijing Advanced Innovation Center for Big Data and Brain Computing; Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China","keywords":"Embedding; Computer science; Probabilistic logic; Encoding (memory); Scalar (mathematics); Artificial intelligence; Transformer; Heuristic; Pattern recognition (psychology); Algorithm; Mathematics","score_opus":0.08892781869382743,"score_gpt":0.3121096265494795,"score_spread":0.22318180785565206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3173999950","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020289376,0.00032433184,0.97675735,0.00036206772,0.000035311758,0.00002354072,0.00013359933,0.0005269207,0.001547604],"genre_scores_gemma":[0.79161775,0.0008855307,0.20165078,0.00030256668,0.00015324862,0.000117301846,0.00046123503,0.00021171778,0.0045997854],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991026,0.00038810552,0.00005619324,0.00023249329,0.00013981329,0.0000808476],"domain_scores_gemma":[0.9977075,0.0012694622,0.00020891921,0.00047487003,0.00023280164,0.00010638288],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015980698,0.00072969834,0.0006426689,0.0007803533,0.00034159695,0.0010846251,0.001391857,0.0008912083,0.0035067615],"category_scores_gemma":[0.009150921,0.00049953314,0.00065336487,0.0011523498,0.0013054643,0.0051715616,0.0025793642,0.0020211893,0.0005621055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031411205,0.00007714035,0.001603415,0.00017216303,0.000054180167,0.00012273407,0.00046630317,0.35272375,0.0055415262,0.37996024,0.0031910806,0.2557734],"study_design_scores_gemma":[0.0000132700825,0.000048052418,0.00023202054,0.000013321996,0.000010538093,0.000035275392,0.000022471726,0.83455116,0.001036238,0.16292243,0.0011019793,0.0000132257355],"about_ca_topic_score_codex":0.0034404788,"about_ca_topic_score_gemma":0.0030932762,"teacher_disagreement_score":0.0035067615,"about_ca_system_score_codex":0.0012567904,"about_ca_system_score_gemma":0.0007593474,"threshold_uncertainty_score":0.011731267},"labels":[],"label_agreement":null},{"id":"W3174012329","doi":"10.48550/arxiv.2012.04056","title":"Semantics Altering Modifications for Evaluating Comprehension in Machine Reading","year":2020,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Computer science; Semantics (computer science); Artificial intelligence; Natural language processing; Process (computing); Sentence; Domain (mathematical analysis); Reading (process); Comprehension; Machine learning; Programming language; Linguistics","score_opus":0.3421025677992835,"score_gpt":0.37227902784880945,"score_spread":0.030176460049525955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174012329","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.64863735,0.0044370317,0.3123913,0.0013163083,0.0003382668,0.00082815334,0.0028539249,0.014287154,0.014910475],"genre_scores_gemma":[0.91177946,0.00043341532,0.08209822,0.0001732748,0.00006185325,0.00026625,0.0030667698,0.0003268475,0.0017939435],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996855,0.0015772535,0.00027387834,0.0006597098,0.0005059408,0.00012808417],"domain_scores_gemma":[0.9822013,0.013004689,0.0011023628,0.0020274918,0.0013054641,0.0003587772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005779427,0.0019285758,0.0007945956,0.0020091112,0.00046505302,0.0023900033,0.0018011796,0.0026525245,0.0040576393],"category_scores_gemma":[0.033248857,0.0004639474,0.00080635096,0.0012750729,0.0010691278,0.0052012913,0.0018821222,0.003038059,0.001531501],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015696775,0.0008726994,0.03403879,0.0021227985,0.000800289,0.00041074477,0.0019150061,0.20694111,0.065044716,0.004651071,0.010080875,0.67155224],"study_design_scores_gemma":[0.000106312655,0.0014889996,0.020374788,0.00014829157,0.00022810596,0.0003160243,0.00071439776,0.8946229,0.06310638,0.013587045,0.0051836995,0.00012303035],"about_ca_topic_score_codex":0.002512331,"about_ca_topic_score_gemma":0.004179669,"teacher_disagreement_score":0.005779427,"about_ca_system_score_codex":0.0011360608,"about_ca_system_score_gemma":0.0008860687,"threshold_uncertainty_score":0.030564904},"labels":[],"label_agreement":null},{"id":"W3174272772","doi":"10.18653/v1/2021.findings-acl.170","title":"An Evaluation of Disentangled Representation Learning for Texts","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; Vector Institute","keywords":"Computer science; Representation (politics); Natural language processing; ENCODE; Task (project management); Artificial intelligence; Style (visual arts); Context (archaeology); Feature learning; Natural language understanding; Transfer of learning; Natural language; Machine learning","score_opus":0.09035460583850222,"score_gpt":0.3785722752602529,"score_spread":0.28821766942175064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174272772","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5483218,0.01076691,0.3915841,0.003990187,0.0010813179,0.0012859228,0.0099663,0.017944606,0.015058837],"genre_scores_gemma":[0.72317624,0.0016712226,0.22771049,0.00096563075,0.00036525872,0.0007692493,0.037622795,0.00096806505,0.0067510814],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99398696,0.0032666612,0.0003203339,0.001414057,0.0008000383,0.00021196014],"domain_scores_gemma":[0.9829359,0.0116536645,0.00067260535,0.0026796295,0.0014782606,0.0005799606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0105689475,0.0027192067,0.0013227408,0.0021704216,0.00085016503,0.0020239083,0.0029372915,0.003579355,0.0030608794],"category_scores_gemma":[0.03402095,0.0004121348,0.0014340646,0.0018350973,0.0013831418,0.005884066,0.0027248557,0.0040367795,0.001767909],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034231788,0.0024024576,0.0075408393,0.0017065766,0.000881476,0.00035329117,0.00043807825,0.42291427,0.016925763,0.0065094037,0.028971767,0.5079329],"study_design_scores_gemma":[0.00021755119,0.0007740311,0.0018441535,0.00008607981,0.00008159617,0.00013564539,0.00012995172,0.97803706,0.009407699,0.0058233403,0.00342935,0.000033544886],"about_ca_topic_score_codex":0.0042587747,"about_ca_topic_score_gemma":0.005598137,"teacher_disagreement_score":0.0105689475,"about_ca_system_score_codex":0.0016811625,"about_ca_system_score_gemma":0.0012255417,"threshold_uncertainty_score":0.055894613},"labels":[],"label_agreement":null},{"id":"W3174432697","doi":"10.18653/v1/2021.acl-short.51","title":"Exploring Listwise Evidence Reasoning with T5 for Fact Verification","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Computational linguistics; Volume (thermodynamics); Natural language processing; Joint (building); Artificial intelligence; Cognitive science; Psychology; Engineering","score_opus":0.2909852846843825,"score_gpt":0.3030645874900807,"score_spread":0.012079302805698189,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174432697","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01213997,0.000903358,0.97031736,0.002642698,0.00017993856,0.00028753385,0.0010532477,0.0069256783,0.005550244],"genre_scores_gemma":[0.24450077,0.00046482275,0.7475821,0.00080924435,0.00020245575,0.00020155993,0.002757542,0.0008762899,0.0026052208],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9895573,0.003998149,0.0009145586,0.0021278458,0.0025400168,0.00086211786],"domain_scores_gemma":[0.9448861,0.04353536,0.0016518754,0.0059655244,0.0031469355,0.000814282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010512958,0.0014349266,0.0016586542,0.0034719147,0.00253487,0.0071295556,0.0058753956,0.0028718642,0.01865436],"category_scores_gemma":[0.054117784,0.001819193,0.0061478964,0.0026906175,0.00323686,0.016841345,0.009023941,0.005201864,0.0045077056],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018824108,0.0006927579,0.0060251537,0.002014621,0.0009463456,0.0018763135,0.002006476,0.076763555,0.009465312,0.48937306,0.04195459,0.36699933],"study_design_scores_gemma":[0.00016593472,0.00009016846,0.0002545831,0.00017000249,0.00025174647,0.00027820922,0.00036740082,0.37668517,0.0089345435,0.59817064,0.014566226,0.00006540692],"about_ca_topic_score_codex":0.011940265,"about_ca_topic_score_gemma":0.023362676,"teacher_disagreement_score":0.01865436,"about_ca_system_score_codex":0.0026822607,"about_ca_system_score_gemma":0.0056604343,"threshold_uncertainty_score":0.06240499},"labels":[],"label_agreement":null},{"id":"W3174593637","doi":"10.21428/594757db.4ea59c2e","title":"Enhancing Pretrained Models with Domain Knowledge","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; Ford Motor Company","keywords":"Computer science; Variety (cybernetics); Domain (mathematical analysis); Domain knowledge; Artificial intelligence; Natural language processing; Language model; Software; Machine learning; Programming language","score_opus":0.02137786640740907,"score_gpt":0.23847327762125453,"score_spread":0.21709541121384546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174593637","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14191139,0.005886023,0.81625336,0.0027917668,0.0008742332,0.00033080898,0.0028635017,0.017437182,0.011651759],"genre_scores_gemma":[0.7472985,0.0031964083,0.21486245,0.0025995504,0.0005868941,0.0006366591,0.016710645,0.0012908814,0.012817963],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99917775,0.00024847063,0.00004417582,0.0003359074,0.00010730022,0.00008639561],"domain_scores_gemma":[0.99561524,0.0030094618,0.00015057616,0.0006184394,0.00050464965,0.00010155248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015483695,0.0024520783,0.001143309,0.0016726803,0.00045872485,0.0014536524,0.00200324,0.0020685606,0.002858151],"category_scores_gemma":[0.0070249154,0.0009127233,0.0013252982,0.0014327407,0.0007515337,0.0033068797,0.0014382196,0.004554211,0.0029974084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028910555,0.00065203995,0.006337454,0.00042139133,0.00040768375,0.0003688394,0.00032757904,0.49709156,0.013865684,0.003401054,0.027621645,0.44921595],"study_design_scores_gemma":[0.00001710211,0.00005838024,0.0008505417,0.000055610155,0.00006588451,0.000062133346,0.000051318053,0.9859483,0.0043065655,0.0040762876,0.004484024,0.000023885717],"about_ca_topic_score_codex":0.008116751,"about_ca_topic_score_gemma":0.018669952,"teacher_disagreement_score":0.008116751,"about_ca_system_score_codex":0.00089956773,"about_ca_system_score_gemma":0.0013749688,"threshold_uncertainty_score":0.01613903},"labels":[],"label_agreement":null},{"id":"W3174693043","doi":"10.18653/v1/2021.acl-long.396","title":"Reasoning over Entity-Action-Location Graph for Procedural Text Understanding","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Natural language processing; Computational linguistics; Artificial intelligence; Action (physics); Graph; Information retrieval; Linguistics; Theoretical computer science; Philosophy","score_opus":0.07962312726237912,"score_gpt":0.3026559286333474,"score_spread":0.22303280137096826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174693043","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014257337,0.00064357516,0.9621236,0.0010784586,0.00007600591,0.00034597536,0.0066488967,0.012655054,0.002171153],"genre_scores_gemma":[0.33532006,0.0010673649,0.6330334,0.00049135904,0.000106274085,0.00045462075,0.025447797,0.0008660151,0.0032130918],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99836105,0.00038602194,0.00017706015,0.00060590013,0.00033360106,0.00013630236],"domain_scores_gemma":[0.99638236,0.0023487601,0.00027945507,0.0005553662,0.00028378493,0.00015023758],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014627791,0.0023216887,0.0014563972,0.0052403305,0.0017377407,0.0030566568,0.0036931583,0.0026286133,0.007690551],"category_scores_gemma":[0.006676795,0.0011195564,0.0049252138,0.0031625077,0.0013412623,0.011783555,0.0036210571,0.0025884002,0.0015911857],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008743102,0.0007219566,0.0081294,0.0028659734,0.0011566954,0.0037372997,0.0030802202,0.2604169,0.015622553,0.28824505,0.075537324,0.33961236],"study_design_scores_gemma":[0.000089409645,0.00004522616,0.00061905716,0.000117609736,0.00035057944,0.0002720073,0.000511357,0.73241633,0.0076402654,0.240908,0.016973408,0.00005675543],"about_ca_topic_score_codex":0.03024841,"about_ca_topic_score_gemma":0.052292563,"teacher_disagreement_score":0.03024841,"about_ca_system_score_codex":0.0021417923,"about_ca_system_score_gemma":0.0024461509,"threshold_uncertainty_score":0.060144663},"labels":[],"label_agreement":null},{"id":"W3174717610","doi":"10.1007/978-3-030-80599-9_28","title":"Using Document Embeddings for Background Linking of News Articles","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Embedding; Pooling; Information retrieval; Variety (cybernetics); Security token; Natural language processing; Artificial intelligence","score_opus":0.06072423983579282,"score_gpt":0.3040035160247984,"score_spread":0.24327927618900558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174717610","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12573273,0.0101844845,0.7878379,0.0013890933,0.0040656356,0.00047945938,0.013045381,0.03602471,0.021240622],"genre_scores_gemma":[0.45826846,0.0048024217,0.4464044,0.0004246914,0.0027358492,0.0003726273,0.05390292,0.003520703,0.029567838],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988802,0.00026222784,0.00009667877,0.00043865343,0.00019224413,0.000130001],"domain_scores_gemma":[0.9963003,0.0018993897,0.0002672757,0.00054961524,0.00075716485,0.00022641705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015533632,0.0022587948,0.0010716307,0.006961254,0.00088596996,0.002929856,0.0010714357,0.001728956,0.007684884],"category_scores_gemma":[0.006627091,0.0006969051,0.0012546752,0.0048799333,0.00028225087,0.0050696842,0.0020715138,0.0020614474,0.009282964],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009823543,0.0006412656,0.004914255,0.0005483353,0.00026320136,0.00031035184,0.0003602342,0.010614953,0.014907,0.0039843065,0.055327356,0.90714633],"study_design_scores_gemma":[0.00016773898,0.0004401254,0.0075818095,0.00029551936,0.0006783026,0.0006499334,0.00070572615,0.8815923,0.029094804,0.025887858,0.052783094,0.00012280593],"about_ca_topic_score_codex":0.0030420842,"about_ca_topic_score_gemma":0.004965898,"teacher_disagreement_score":0.007684884,"about_ca_system_score_codex":0.0005967126,"about_ca_system_score_gemma":0.0007412162,"threshold_uncertainty_score":0.025708497},"labels":[],"label_agreement":null},{"id":"W3174726724","doi":"10.18653/v1/2021.acl-long.163","title":"Optimizing Deeper Transformers on Small Datasets","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University; Canadian Institute for Advanced Research; University of Alberta; University of Waterloo","funders":"","keywords":"Computer science; Computational linguistics; Natural language processing; Transformer; Artificial intelligence; Library science; Cognitive science; Engineering; Psychology; Electrical engineering","score_opus":0.039229325929457916,"score_gpt":0.2517262054989651,"score_spread":0.21249687956950722,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174726724","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25438166,0.0052562393,0.650514,0.003146841,0.0008041891,0.000723149,0.009363484,0.05892189,0.0168886],"genre_scores_gemma":[0.6183723,0.0008286446,0.344199,0.00083481596,0.000287657,0.00033855642,0.020949,0.004556961,0.009633103],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99694353,0.00085051893,0.00026308023,0.0009592454,0.000517633,0.00046587703],"domain_scores_gemma":[0.9910102,0.0051392163,0.00021930061,0.0026028575,0.00077240204,0.00025609465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037295918,0.0022388115,0.002996166,0.0020908078,0.0013487481,0.002988059,0.0029642363,0.0015806601,0.014226249],"category_scores_gemma":[0.017975112,0.0013644574,0.0025280036,0.0034312261,0.0012987037,0.0132950265,0.0039092847,0.0032733106,0.0035693725],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001775751,0.0008297619,0.008462082,0.00085897226,0.00045664114,0.00029860903,0.00051789737,0.15523928,0.011342668,0.05194651,0.10080196,0.66746986],"study_design_scores_gemma":[0.00030921798,0.00019502228,0.0009575213,0.00006141548,0.00016998258,0.00012412033,0.00028976667,0.86147106,0.0048205154,0.12539685,0.0061816094,0.000023026281],"about_ca_topic_score_codex":0.0092453,"about_ca_topic_score_gemma":0.032092508,"teacher_disagreement_score":0.014226249,"about_ca_system_score_codex":0.0022966922,"about_ca_system_score_gemma":0.0041325847,"threshold_uncertainty_score":0.047591567},"labels":[],"label_agreement":null},{"id":"W3174851372","doi":"10.18653/v1/2021.findings-acl.236","title":"GrantRel: Grant Information Extraction via Joint Entity and Relation Extraction","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; Higher Education Discipline Innovation Project; Canadian Institutes of Health Research; Science and Technology Commission of Shanghai Municipality; National Natural Science Foundation of China","keywords":"Relationship extraction; Joint (building); Computer science; Extraction (chemistry); Information extraction; Relation (database); Information retrieval; Data mining; Artificial intelligence; Engineering; Structural engineering; Chromatography","score_opus":0.02143575687510175,"score_gpt":0.24138614835610275,"score_spread":0.219950391481001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3174851372","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016963936,0.0026770888,0.86104316,0.0020053769,0.0005544759,0.0012623031,0.051945657,0.05484876,0.00869918],"genre_scores_gemma":[0.087671585,0.0017787471,0.77390295,0.0006270188,0.000630958,0.001360107,0.11907684,0.0019608012,0.012991042],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9958418,0.00078189885,0.00070400746,0.0012225015,0.0011369262,0.00031285067],"domain_scores_gemma":[0.98988473,0.004447657,0.0011639656,0.0019705251,0.00216057,0.0003725958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060952245,0.0022394431,0.0015016233,0.010802855,0.0011103426,0.0026551157,0.0021016733,0.0018796533,0.008849998],"category_scores_gemma":[0.015947705,0.0009105895,0.0020020152,0.0077757672,0.0005630315,0.007210221,0.0040470413,0.001929842,0.009520538],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038757746,0.00040275627,0.010243889,0.0017700206,0.00038477153,0.001227825,0.0011042476,0.0051577836,0.02758429,0.0228463,0.23797368,0.69091684],"study_design_scores_gemma":[0.00025476137,0.00046692332,0.016479556,0.0005138318,0.0008641016,0.0025369616,0.0012322069,0.24825117,0.093683474,0.08693236,0.5484329,0.00035180443],"about_ca_topic_score_codex":0.0057525127,"about_ca_topic_score_gemma":0.009101983,"teacher_disagreement_score":0.010802855,"about_ca_system_score_codex":0.00093976606,"about_ca_system_score_gemma":0.004648365,"threshold_uncertainty_score":0.032235026},"labels":[],"label_agreement":null},{"id":"W3175052694","doi":"10.18653/v1/2021.acl-long.377","title":"Turn the Combination Lock: Learnable Textual Backdoor Attacks via Word Substitution","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"National Key Research and Development Program of China","keywords":"Backdoor; Substitution (logic); Computer science; Lock (firearm); Word (group theory); Computational linguistics; Natural language processing; Artificial intelligence; Linguistics; History; Programming language; Philosophy; Computer security","score_opus":0.021426203902997237,"score_gpt":0.24675431153586064,"score_spread":0.2253281076328634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3175052694","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08741924,0.0028110566,0.8647347,0.0023448137,0.0010607743,0.00036138593,0.0009948389,0.02279363,0.017479517],"genre_scores_gemma":[0.88255095,0.00051861635,0.10264872,0.0011210587,0.00037069758,0.0002839174,0.00092959777,0.0011542218,0.010422321],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964318,0.0010418665,0.00026052602,0.0006888222,0.0008964952,0.0006805167],"domain_scores_gemma":[0.9935035,0.002776259,0.00041057612,0.0026876251,0.00040562442,0.00021635924],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019121653,0.0012185711,0.0017852886,0.0011075246,0.001104886,0.0022458052,0.0021799516,0.0018980128,0.009940633],"category_scores_gemma":[0.0076199733,0.0007792272,0.001140633,0.0011777214,0.0020708297,0.008061816,0.0065275505,0.0030245865,0.004258008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039312993,0.0009051402,0.0032893524,0.0008402748,0.00056076335,0.0009558066,0.0007778816,0.06480153,0.051186655,0.17138308,0.07930561,0.6220627],"study_design_scores_gemma":[0.00041763944,0.00060398836,0.00047951733,0.000105100684,0.00024626844,0.000587126,0.00018687594,0.6087631,0.041538008,0.3261041,0.020860994,0.00010730421],"about_ca_topic_score_codex":0.0005727278,"about_ca_topic_score_gemma":0.0011103203,"teacher_disagreement_score":0.009940633,"about_ca_system_score_codex":0.0007702814,"about_ca_system_score_gemma":0.0012996813,"threshold_uncertainty_score":0.033254683},"labels":[],"label_agreement":null},{"id":"W3175560839","doi":"10.18653/v1/2021.acl-long.325","title":"How is BERT surprised? Layerwise detection of linguistic anomalies","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Volume (thermodynamics); Linguistics; Joint (building); Computational linguistics; Natural language processing; Artificial intelligence; Philosophy; Engineering; Physics; Structural engineering","score_opus":0.01879292536861095,"score_gpt":0.2284517515391103,"score_spread":0.20965882617049936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3175560839","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67123795,0.0037170663,0.30169797,0.007925926,0.00123208,0.00004572017,0.0021196227,0.0048670215,0.0071566305],"genre_scores_gemma":[0.9734292,0.00037892858,0.022109589,0.00045375086,0.00026778964,0.000017599168,0.0013241065,0.000316196,0.0017027411],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990994,0.00022781966,0.000045627858,0.00029283832,0.00016993466,0.0001644577],"domain_scores_gemma":[0.99695873,0.0011962207,0.00037297807,0.00049517164,0.00074895856,0.00022795855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019544582,0.00068470626,0.0009283808,0.0008434701,0.000496686,0.001574808,0.0012103671,0.0009466661,0.0014716078],"category_scores_gemma":[0.011785748,0.00033884036,0.00035637204,0.0008128161,0.00058612204,0.0028382675,0.0012311555,0.0017042025,0.0011461409],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038922555,0.00022960924,0.11927996,0.00031359628,0.00044036127,0.0011660119,0.0009216905,0.026656328,0.0848094,0.011599662,0.060229596,0.6904616],"study_design_scores_gemma":[0.000050826147,0.0002262714,0.04245364,0.00007005658,0.00025649593,0.0006203429,0.0007832038,0.86492515,0.027259706,0.053362146,0.009909981,0.00008216089],"about_ca_topic_score_codex":0.0026914119,"about_ca_topic_score_gemma":0.0033702368,"teacher_disagreement_score":0.0026914119,"about_ca_system_score_codex":0.0004901467,"about_ca_system_score_gemma":0.00059455616,"threshold_uncertainty_score":0.01033622},"labels":[],"label_agreement":null},{"id":"W3175902990","doi":"","title":"H2oloo at TREC 2020: When all you got is a hammer... Deep Learning, Health Misinformation, and Precision Medicine.","year":2020,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Misinformation; Hammer; Computer science; Artificial intelligence; Deep learning; Data science; Engineering; Computer security; Mechanical engineering","score_opus":0.056946370647176676,"score_gpt":0.28683623493138527,"score_spread":0.2298898642842086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3175902990","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017278662,0.036881328,0.019241523,0.17164502,0.06058912,0.0014931319,0.5981041,0.0111066755,0.08366036],"genre_scores_gemma":[0.08486571,0.007941765,0.029418724,0.04291128,0.015510723,0.0011276633,0.6851285,0.0029143235,0.13018134],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969049,0.001304041,0.000121543955,0.0002847994,0.001006792,0.0003778348],"domain_scores_gemma":[0.9836233,0.006292259,0.0006661137,0.0010091203,0.0048691765,0.003539946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0114950575,0.0023484242,0.0013452793,0.0026670913,0.0020323028,0.0050025433,0.0025201712,0.004764943,0.046827268],"category_scores_gemma":[0.017502062,0.00048533388,0.0010121892,0.0017601597,0.0017939064,0.0045013507,0.0030821783,0.0042768237,0.018636195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012028656,0.000027326958,0.00026387352,0.00022146778,0.000029315896,0.000012092934,0.000016130103,0.00040368608,0.00022455072,0.00028453165,0.9913121,0.007084661],"study_design_scores_gemma":[0.001433197,0.00051513966,0.011870504,0.0012007034,0.00029002898,0.00015748719,0.00061670406,0.016127871,0.0051880353,0.021747675,0.9405427,0.00030999462],"about_ca_topic_score_codex":0.0801805,"about_ca_topic_score_gemma":0.20271587,"teacher_disagreement_score":0.0801805,"about_ca_system_score_codex":0.004015107,"about_ca_system_score_gemma":0.0076537393,"threshold_uncertainty_score":0.15942758},"labels":[],"label_agreement":null},{"id":"W3176125464","doi":"10.18653/v1/2021.acl-short.137","title":"Bringing Structure into Summaries: a Faceted Summarization Dataset for Long Scientific Documents","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Computer science; Zhàng; Volume (thermodynamics); Joint (building); Computational linguistics; Natural language processing; Information retrieval; Library science; History; China; Engineering; Archaeology","score_opus":0.017831990907906967,"score_gpt":0.2763486229645114,"score_spread":0.2585166320566044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176125464","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059493642,0.009797599,0.015581802,0.0021440627,0.0006576501,0.0006273824,0.8974096,0.009436242,0.0048520276],"genre_scores_gemma":[0.030108398,0.00092323247,0.026069893,0.00014882095,0.00017216084,0.00042171543,0.94006866,0.00017851379,0.0019086259],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985065,0.00034944498,0.00026706688,0.00042238308,0.00032470634,0.00012978824],"domain_scores_gemma":[0.995669,0.0012567098,0.000427914,0.0008072162,0.0013100164,0.000529142],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0015980479,0.0014403332,0.0010552839,0.006490723,0.0013948486,0.0016379397,0.0015183439,0.0018910294,0.003683512],"category_scores_gemma":[0.0069634034,0.000346017,0.0010475542,0.0057308963,0.00037377086,0.0020487995,0.0019484125,0.000963216,0.004671366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001279033,0.00061436865,0.016063983,0.004494131,0.0006510706,0.0007039403,0.001669301,0.003344842,0.017474145,0.0024076523,0.743755,0.20754254],"study_design_scores_gemma":[0.0012192976,0.0013619708,0.07597423,0.0007719376,0.0010772652,0.001343077,0.0048803124,0.045351263,0.025246937,0.0147849135,0.8275497,0.00043905078],"about_ca_topic_score_codex":0.0076167,"about_ca_topic_score_gemma":0.0257215,"teacher_disagreement_score":0.99840194,"about_ca_system_score_codex":0.00087458355,"about_ca_system_score_gemma":0.0019085304,"threshold_uncertainty_score":0.015144706},"labels":[],"label_agreement":null},{"id":"W3176158018","doi":"10.48550/arxiv.2104.01940","title":"What's the best place for an AI conference, Vancouver or ______: Why completing comparative questions is difficult","year":2021,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Competence (human resources); sort; Natural language processing; Benchmark (surveying); Language model; Set (abstract data type); Task (project management); Deep learning; Machine learning; Cognitive science; Information retrieval; Psychology","score_opus":0.19003925533081464,"score_gpt":0.25084525373212957,"score_spread":0.060805998401314926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176158018","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20685199,0.008965326,0.028013293,0.08861245,0.0052768486,0.00040656122,0.022294668,0.006677674,0.6329013],"genre_scores_gemma":[0.72575736,0.0045999265,0.040845968,0.0046008085,0.00071358343,0.00020171981,0.024805414,0.00114533,0.19732982],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993818,0.000213152,0.000029081277,0.00015423336,0.00012464529,0.00009718076],"domain_scores_gemma":[0.9976567,0.00066793297,0.00012631781,0.00020329379,0.0007211544,0.0006245983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012650226,0.0005278846,0.0003784317,0.00072055403,0.0029897175,0.0043262816,0.00078060734,0.0014738974,0.047739103],"category_scores_gemma":[0.0077515896,0.00024405963,0.00032064688,0.0013764395,0.00091959577,0.0034259702,0.00089846057,0.0020316879,0.01043873],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005182133,0.0001723944,0.01747566,0.00068985444,0.00008706671,0.0006906069,0.0031444838,0.005383907,0.0037277266,0.028661245,0.62719107,0.3122578],"study_design_scores_gemma":[0.000064700434,0.00009269969,0.031670008,0.00047903587,0.00005406393,0.0004226626,0.008216365,0.014671966,0.004436414,0.031686652,0.90807265,0.00013275679],"about_ca_topic_score_codex":0.17436738,"about_ca_topic_score_gemma":0.40023807,"teacher_disagreement_score":0.17436738,"about_ca_system_score_codex":0.00356364,"about_ca_system_score_gemma":0.0037399887,"threshold_uncertainty_score":0.3467049},"labels":[],"label_agreement":null},{"id":"W3176169354","doi":"10.18653/v1/2021.acl-long.551","title":"ARBERT &amp; MARBERT: Deep Bidirectional Transformers for Arabic","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":352,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Arabic; Transformer; Computer science; Natural language processing; Linguistics; Artificial intelligence; Speech recognition; Electrical engineering; Engineering; Philosophy; Voltage","score_opus":0.028567093499341254,"score_gpt":0.2625137650090602,"score_spread":0.23394667150971896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176169354","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016231757,0.0022608123,0.8129079,0.0017969914,0.00093197305,0.00015035996,0.007943547,0.11354484,0.044231787],"genre_scores_gemma":[0.3332583,0.001724971,0.5817866,0.000833377,0.0002358029,0.00023928087,0.01438654,0.012096773,0.05543839],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99979514,0.00004894977,0.000015368072,0.000059477374,0.000050797094,0.000030290012],"domain_scores_gemma":[0.9996636,0.0001275425,0.000015890966,0.00008676078,0.000080720616,0.000025640347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004288811,0.0012777895,0.0005017632,0.0010251572,0.0008145044,0.002055159,0.0012746801,0.00086862675,0.048265513],"category_scores_gemma":[0.0015774004,0.0007814677,0.000946492,0.000791222,0.0005474907,0.006130095,0.0020376435,0.0017542295,0.018337036],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006389768,0.00010254363,0.0007182117,0.00047424628,0.00008044076,0.00046166798,0.00050391833,0.01285183,0.013720843,0.095038906,0.28716657,0.5882418],"study_design_scores_gemma":[0.00020319836,0.0001132504,0.0005828761,0.0001980522,0.000102923885,0.0004963652,0.00044402032,0.31940007,0.05625177,0.28829715,0.33381534,0.000094979834],"about_ca_topic_score_codex":0.005249097,"about_ca_topic_score_gemma":0.013400765,"teacher_disagreement_score":0.048265513,"about_ca_system_score_codex":0.0007569564,"about_ca_system_score_gemma":0.00075471884,"threshold_uncertainty_score":0.16146421},"labels":[],"label_agreement":null},{"id":"W3176527752","doi":"","title":"Spotify at the TREC 2020 Podcasts Track: Segment Retrieval.","year":2020,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Track (disk drive); Computer science; Information retrieval; World Wide Web; Operating system","score_opus":0.054511271836274806,"score_gpt":0.2668095807724929,"score_spread":0.2122983089362181,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176527752","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034919076,0.015380205,0.07520244,0.0288929,0.03860525,0.003982918,0.56377023,0.09874078,0.14050621],"genre_scores_gemma":[0.03192368,0.0025518541,0.04111877,0.0035616904,0.005238068,0.0009285126,0.675215,0.0056847143,0.23377772],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997855,0.00047744752,0.000081100654,0.00031481875,0.00092077826,0.00035078224],"domain_scores_gemma":[0.99247736,0.001319226,0.00018120407,0.00084283756,0.0033450776,0.0018342716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046977256,0.0026059018,0.0022001823,0.003562521,0.0026171284,0.005394547,0.0024253225,0.0027788312,0.091113925],"category_scores_gemma":[0.007243303,0.0005641526,0.0009149056,0.0033326454,0.0006609594,0.0052162865,0.0025776238,0.0031588087,0.06771435],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014283705,0.00007740664,0.00014546225,0.00014337394,0.00001658794,0.000021466798,0.000035711782,0.00017814066,0.0027204596,0.0002875172,0.9815672,0.014663797],"study_design_scores_gemma":[0.0006848253,0.0008269942,0.0078263115,0.00017543275,0.0001547846,0.00019691243,0.0006182693,0.017224012,0.021025648,0.0044549797,0.9466633,0.00014853018],"about_ca_topic_score_codex":0.06878214,"about_ca_topic_score_gemma":0.1561614,"teacher_disagreement_score":0.091113925,"about_ca_system_score_codex":0.0021971222,"about_ca_system_score_gemma":0.0046306215,"threshold_uncertainty_score":0.30480647},"labels":[],"label_agreement":null},{"id":"W3176581245","doi":"10.18653/v1/2021.sigdial-1.49","title":"A Brief Study on the Effects of Training Generative Dialogue Models with a Semantic loss","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Utterance; Computer science; Artificial intelligence; Language model; Security token; Natural language processing; Generative grammar; Training set; Set (abstract data type); Semantic similarity; Generative model; Task (project management); Context (archaeology); Principle of maximum entropy; Machine learning","score_opus":0.05970064700788191,"score_gpt":0.26324544836167996,"score_spread":0.20354480135379804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176581245","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2076923,0.0138112465,0.7551509,0.0037802088,0.00065869035,0.00048319393,0.0007144965,0.002749652,0.014959383],"genre_scores_gemma":[0.7665795,0.0027795946,0.22093411,0.0009934996,0.0004245617,0.00042523234,0.001034815,0.0008679563,0.0059607965],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949656,0.0034473622,0.000247131,0.00061483314,0.00045526194,0.00026974766],"domain_scores_gemma":[0.95193356,0.04315516,0.000694448,0.0026343612,0.0012155104,0.0003669344],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008114505,0.0021706107,0.0015550628,0.0006679775,0.0007979129,0.0020518457,0.0014680044,0.002162036,0.0050766976],"category_scores_gemma":[0.04585599,0.00081216695,0.0013214722,0.0008044898,0.0013768636,0.0037541774,0.0027238103,0.005142096,0.000825885],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021856811,0.0010375092,0.0063261064,0.0011069402,0.00052557,0.00027835296,0.0004863686,0.67411685,0.021283362,0.016856138,0.005569738,0.27022734],"study_design_scores_gemma":[0.00008478059,0.0007701694,0.0017124059,0.00007385511,0.00008001781,0.00013279995,0.00012267916,0.97554433,0.00952759,0.009625912,0.0022908484,0.000034615823],"about_ca_topic_score_codex":0.0047997315,"about_ca_topic_score_gemma":0.005025733,"teacher_disagreement_score":0.008114505,"about_ca_system_score_codex":0.001250001,"about_ca_system_score_gemma":0.0006881308,"threshold_uncertainty_score":0.042914093},"labels":[],"label_agreement":null},{"id":"W3176607063","doi":"10.18653/v1/2021.acl-long.559","title":"StructFormer: Joint Unsupervised Induction of Dependency and Constituency Structure from Masked Language Modeling","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Joint (building); Dependency (UML); Computer science; Natural language processing; Linguistics; Artificial intelligence; Volume (thermodynamics); Computational linguistics; Cognitive science; Philosophy; Engineering; Psychology; Physics; Structural engineering","score_opus":0.02166416538754473,"score_gpt":0.2252969243235708,"score_spread":0.20363275893602606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176607063","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01144912,0.0004370003,0.96087015,0.0003049273,0.00017182696,0.00009754097,0.00209447,0.02280535,0.0017697078],"genre_scores_gemma":[0.19293173,0.0005097256,0.76397306,0.0003971129,0.00031805047,0.0005917846,0.026700668,0.0062683243,0.008309595],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986916,0.00053070736,0.000059304057,0.00042091415,0.00016981816,0.00012771801],"domain_scores_gemma":[0.99709034,0.0017330745,0.00014711956,0.0005310014,0.000342679,0.00015569018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022171934,0.0022511971,0.0013822687,0.0021081425,0.0011713688,0.0017312783,0.002934363,0.001308022,0.007273285],"category_scores_gemma":[0.004403973,0.0015695317,0.0025233282,0.0020019829,0.00086082274,0.003949284,0.0032275717,0.002778633,0.007640688],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009722116,0.00046688234,0.005110839,0.00064016343,0.0006829031,0.0007241608,0.00084503746,0.052433677,0.04168963,0.03558474,0.114150606,0.74669915],"study_design_scores_gemma":[0.00013495192,0.00011677456,0.0012474385,0.000034103785,0.00014454126,0.00016241775,0.00012504327,0.9196001,0.01389898,0.050711896,0.013770131,0.000053608797],"about_ca_topic_score_codex":0.005310724,"about_ca_topic_score_gemma":0.014900826,"teacher_disagreement_score":0.007273285,"about_ca_system_score_codex":0.0007297703,"about_ca_system_score_gemma":0.0021886097,"threshold_uncertainty_score":0.02433157},"labels":[],"label_agreement":null},{"id":"W3176616156","doi":"10.1609/aaai.v35i16.17722","title":"Neural Sentence Ordering Based on Constraint Graphs","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Science Foundation of Shandong Province; National Natural Science Foundation of China","keywords":"Sentence; Computer science; Constraint (computer-aided design); Graph; Artificial intelligence; Benchmark (surveying); Natural language processing; Code (set theory); Theoretical computer science; Mathematics; Programming language","score_opus":0.08688291472376188,"score_gpt":0.2864605725247814,"score_spread":0.19957765780101955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176616156","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10904259,0.0017271398,0.8505158,0.001122314,0.00038856623,0.00056148705,0.007815985,0.014399312,0.014426838],"genre_scores_gemma":[0.5865705,0.0010726596,0.36957315,0.0006086677,0.00027940425,0.00045035823,0.024842272,0.0009682414,0.01563482],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994338,0.00012528093,0.000036201574,0.00021881639,0.00013285638,0.000053040454],"domain_scores_gemma":[0.99917275,0.00032777915,0.00011721539,0.00011625697,0.0002054087,0.000060562666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00044193363,0.0012099693,0.000565029,0.0022328352,0.00046577773,0.00080351107,0.0010695817,0.0006723336,0.0050572506],"category_scores_gemma":[0.0029422867,0.0003655773,0.00090790004,0.0020112358,0.0004413783,0.002606192,0.00079670944,0.0013930743,0.001502694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057769,0.00038715245,0.0037867003,0.0006649546,0.0001893527,0.00044422862,0.00045732094,0.112018295,0.031821188,0.027501399,0.048319846,0.77383184],"study_design_scores_gemma":[0.0000709441,0.0001528624,0.0027470673,0.00005782392,0.000105421896,0.00019344139,0.00016492985,0.90830684,0.012288839,0.062409814,0.013454599,0.000047468366],"about_ca_topic_score_codex":0.011119196,"about_ca_topic_score_gemma":0.024723494,"teacher_disagreement_score":0.011119196,"about_ca_system_score_codex":0.0010574491,"about_ca_system_score_gemma":0.0015911813,"threshold_uncertainty_score":0.022108912},"labels":[],"label_agreement":null},{"id":"W3176692111","doi":"10.18653/v1/2021.findings-acl.119","title":"LICHEE: Improving Language Model Pre-training with Multi-grained Tokenization","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Lexical analysis; Language model; Benchmark (surveying); Artificial intelligence; Natural language processing; Inference; Variety (cybernetics); Natural language understanding; Representation (politics); Natural language; Machine learning","score_opus":0.03371922272302266,"score_gpt":0.2704040480036245,"score_spread":0.23668482528060183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176692111","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030984111,0.001412377,0.932345,0.00044868165,0.0003925011,0.00018597604,0.0010626647,0.030635243,0.0025335338],"genre_scores_gemma":[0.3407575,0.00089796865,0.6233751,0.0013866372,0.0002916917,0.00082093745,0.013889852,0.0034098702,0.015170458],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985241,0.00047024174,0.000097050695,0.00049515,0.00021889011,0.00019462864],"domain_scores_gemma":[0.99692744,0.0016609324,0.000112917965,0.0006994924,0.00045974695,0.00013949454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022045537,0.002294614,0.0018788769,0.0019008268,0.0008585251,0.0012685655,0.00314673,0.0016157996,0.005689275],"category_scores_gemma":[0.006913162,0.0010736049,0.0015817041,0.0017513378,0.00081779796,0.0053218137,0.0029312493,0.005041662,0.005197589],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048221496,0.0005163775,0.0030125277,0.00035145343,0.00030249372,0.00033548783,0.00041915188,0.1639594,0.019459335,0.005935372,0.037810616,0.7674156],"study_design_scores_gemma":[0.000055464123,0.00011584127,0.0004919382,0.00002438743,0.000047749178,0.000098497236,0.000089269335,0.97723764,0.010232404,0.005148555,0.006410383,0.00004783605],"about_ca_topic_score_codex":0.012960012,"about_ca_topic_score_gemma":0.0296892,"teacher_disagreement_score":0.012960012,"about_ca_system_score_codex":0.0009651424,"about_ca_system_score_gemma":0.002803972,"threshold_uncertainty_score":0.025769174},"labels":[],"label_agreement":null},{"id":"W3176765152","doi":"10.18653/v1/2021.acl-long.465","title":"Guiding the Growth: Difficulty-Controllable Question Generation through Step-by-Step Rewriting","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Rewriting; Computer science; Natural language generation; Programming language; Computational linguistics; Linguistics; Natural language processing; Library science; Natural language; Philosophy","score_opus":0.0411296921574314,"score_gpt":0.2590174505973326,"score_spread":0.2178877584399012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176765152","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02952191,0.00070229097,0.94885826,0.00078505906,0.00021981113,0.00044316935,0.0007742641,0.012318029,0.006377201],"genre_scores_gemma":[0.40977508,0.00036913692,0.5758458,0.00046812504,0.00013389859,0.000498102,0.0030936783,0.0031502924,0.0066658724],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99759537,0.0010947295,0.00015184231,0.00055463804,0.00043894586,0.00016441933],"domain_scores_gemma":[0.99089843,0.0062418035,0.00022084397,0.0013643816,0.0010285092,0.00024603962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022968745,0.0012538509,0.0012347414,0.00092806784,0.00070696085,0.0015945619,0.0031635435,0.0012655455,0.008512961],"category_scores_gemma":[0.015540758,0.00079785095,0.0015126113,0.000841058,0.0010549302,0.0038630918,0.0037740113,0.0018852593,0.0037830665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076730386,0.00069192715,0.0048926114,0.0014476366,0.0002853753,0.0011673014,0.0044570277,0.07055412,0.10032236,0.10162262,0.077591,0.6362007],"study_design_scores_gemma":[0.00011630985,0.000097846365,0.0004546037,0.000048730755,0.0001305429,0.00021314863,0.00046525005,0.8460053,0.026788028,0.10569865,0.019932644,0.000048878203],"about_ca_topic_score_codex":0.0022798781,"about_ca_topic_score_gemma":0.0047846884,"teacher_disagreement_score":0.008512961,"about_ca_system_score_codex":0.0006712033,"about_ca_system_score_gemma":0.0010250991,"threshold_uncertainty_score":0.028478742},"labels":[],"label_agreement":null},{"id":"W3176889315","doi":"","title":"Impact of Tokenization, Pretraining Task, and Transformer Depth on Text Ranking","year":2021,"lang":"en","type":"article","venue":"UvA-DARE (University of Amsterdam)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Canadian Institute of Steel Construction","keywords":"Computer science; Transformer; Lexical analysis; Preprocessor; Artificial intelligence; Natural language processing; Vocabulary; Machine learning; Question answering; Information retrieval; Linguistics","score_opus":0.0165175014853993,"score_gpt":0.22467212070002332,"score_spread":0.20815461921462403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176889315","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7418754,0.0076675722,0.2001847,0.0028703273,0.0011972963,0.0006711889,0.0018311219,0.0243644,0.01933802],"genre_scores_gemma":[0.9172028,0.0011399781,0.06683291,0.00094704813,0.00021025051,0.00020692634,0.0034205276,0.0016549111,0.008384642],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978167,0.0008734432,0.00018816888,0.00047397928,0.00029260296,0.00035513088],"domain_scores_gemma":[0.98737144,0.009069157,0.00045211555,0.0016244036,0.00086352095,0.0006194048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004420233,0.0015029766,0.0010496031,0.00065666065,0.0007462113,0.002145054,0.0016582586,0.0015821258,0.006754209],"category_scores_gemma":[0.03161162,0.0005209697,0.00066037383,0.0009030909,0.0010748426,0.008130583,0.0024178296,0.002686282,0.003685436],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0072121164,0.0014386352,0.010984708,0.0011018411,0.00019740028,0.00035086554,0.00043988132,0.098183334,0.04724128,0.005710559,0.018988635,0.80815077],"study_design_scores_gemma":[0.0010840426,0.00680234,0.012975735,0.00035283258,0.00076772145,0.0008482089,0.0014144583,0.82508415,0.107520394,0.023377046,0.019549293,0.00022369406],"about_ca_topic_score_codex":0.005941891,"about_ca_topic_score_gemma":0.0070697377,"teacher_disagreement_score":0.006754209,"about_ca_system_score_codex":0.0010289052,"about_ca_system_score_gemma":0.002446692,"threshold_uncertainty_score":0.023376703},"labels":[],"label_agreement":null},{"id":"W3176904855","doi":"10.1609/aaai.v35i15.17608","title":"On the Softmax Bottleneck of Recurrent Language Models","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada; Western Canada Research Grid; Compute Canada","keywords":"Perplexity; Softmax function; Bottleneck; Rank (graph theory); Computer science; Monotonic function; Word (group theory); Correlation; Language model; Mathematics; Artificial intelligence; Artificial neural network; Combinatorics","score_opus":0.10595509279220487,"score_gpt":0.2987821668527856,"score_spread":0.1928270740605807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176904855","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04657959,0.001027933,0.94560826,0.0011210921,0.000071683986,0.000040622494,0.00020298081,0.0013047816,0.0040429565],"genre_scores_gemma":[0.88682276,0.001568538,0.10084964,0.000711217,0.0002605008,0.00026126314,0.0007769778,0.00054423045,0.00820482],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973597,0.0014376523,0.00012610632,0.00043011765,0.00039904463,0.00024741777],"domain_scores_gemma":[0.9846042,0.012572356,0.00061916857,0.00092216616,0.0009699545,0.00031220246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0079770405,0.0014743031,0.0016187609,0.0009919953,0.00074653,0.0022423847,0.0020089229,0.0017783873,0.0034191245],"category_scores_gemma":[0.034188684,0.0012533242,0.0010970446,0.0009420485,0.0017594729,0.0062038903,0.0027034201,0.004213871,0.0014086008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049838657,0.0001229758,0.0017142933,0.0002429454,0.00013407345,0.00028243518,0.00031290035,0.7751242,0.0042505106,0.11632879,0.004117387,0.09687114],"study_design_scores_gemma":[0.000007196802,0.000026837593,0.00012213878,0.000014027517,0.000008133074,0.000020924734,0.000010464587,0.96591914,0.00077188737,0.03275678,0.00033236446,0.000009987245],"about_ca_topic_score_codex":0.0064134495,"about_ca_topic_score_gemma":0.004854272,"teacher_disagreement_score":0.0079770405,"about_ca_system_score_codex":0.0017141642,"about_ca_system_score_gemma":0.0015659424,"threshold_uncertainty_score":0.042187095},"labels":[],"label_agreement":null},{"id":"W3177337627","doi":"10.1609/aaai.v35i14.17524","title":"Analogy Training Multilingual Encoders","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"H. Lundbeck A/S; Natural Sciences and Engineering Research Council of Canada; Lundbeckfonden","keywords":"Analogy; Computer science; Natural language processing; Encoder; Sentence; Artificial intelligence; Word (group theory); Encoding (memory); Linguistics","score_opus":0.1549327852651082,"score_gpt":0.3235714892127046,"score_spread":0.16863870394759642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177337627","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.115667626,0.0009706043,0.8475153,0.0006271271,0.00046754064,0.00021404286,0.0019656657,0.014179009,0.018393045],"genre_scores_gemma":[0.6646568,0.00045592684,0.30589476,0.0005907896,0.00018281084,0.00037676812,0.007655439,0.00094841246,0.019238245],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99900913,0.00025353156,0.000058099322,0.00039853374,0.0001708732,0.00010978456],"domain_scores_gemma":[0.9984358,0.0005003202,0.000056443576,0.00042225956,0.00050428155,0.000080877595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001015749,0.0010339017,0.0006891218,0.0007829671,0.0006514318,0.00093922764,0.0015835892,0.0007290929,0.008361525],"category_scores_gemma":[0.0051741092,0.00047908642,0.00071251765,0.00076101866,0.0005236349,0.0027571202,0.0022122746,0.0019500481,0.0033467792],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042951183,0.00039141937,0.0057171383,0.00039968904,0.0001915426,0.00041992485,0.000571743,0.09722338,0.02651792,0.04842462,0.03290292,0.7868102],"study_design_scores_gemma":[0.00009706498,0.00019301871,0.0010507299,0.00004141095,0.000082207196,0.0002309462,0.00020910343,0.90887725,0.025157444,0.042838965,0.021180509,0.000041312258],"about_ca_topic_score_codex":0.0041344627,"about_ca_topic_score_gemma":0.0115187215,"teacher_disagreement_score":0.008361525,"about_ca_system_score_codex":0.00070396037,"about_ca_system_score_gemma":0.0016038578,"threshold_uncertainty_score":0.027972043},"labels":[],"label_agreement":null},{"id":"W3177367299","doi":"10.1609/aaai.v35i16.17642","title":"Encoding Syntactic Knowledge in Transformer Encoder for Intent Detection and Slot Filling","year":2021,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Transformer; Security token; Encoder; Inference; Parsing; ENCODE; Artificial intelligence; F1 score; Benchmark (surveying); Speech recognition; Natural language processing; Pattern recognition (psychology); Voltage","score_opus":0.10334458186460796,"score_gpt":0.30684030738651535,"score_spread":0.2034957255219074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177367299","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054362,0.00093988294,0.9192593,0.0006301166,0.00034619583,0.00017081617,0.0018973151,0.016262673,0.0061316625],"genre_scores_gemma":[0.7286578,0.0007285285,0.2547856,0.00063996087,0.00014865183,0.00025772423,0.0050958595,0.00065417466,0.009031758],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996661,0.00006390001,0.000022084876,0.00012042218,0.000067646324,0.00005991189],"domain_scores_gemma":[0.99903786,0.00044446113,0.000055499975,0.00017598976,0.00023535256,0.000050918716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008568734,0.001278769,0.00077890756,0.0010714929,0.00033074844,0.0010134117,0.0017599377,0.00099521,0.0037246607],"category_scores_gemma":[0.0026806337,0.0005944632,0.0010689521,0.00077837816,0.0006323197,0.0035117152,0.0012100961,0.0018503837,0.0026583096],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083239394,0.0005660716,0.008116238,0.00048465678,0.0001870096,0.00053865876,0.0005998186,0.08744125,0.04385919,0.021501485,0.023911053,0.8119622],"study_design_scores_gemma":[0.000037155678,0.00011222162,0.00093160715,0.000045746874,0.00012457451,0.00021327242,0.00006517791,0.9456095,0.025795275,0.02106118,0.0059620524,0.000042139534],"about_ca_topic_score_codex":0.0063129645,"about_ca_topic_score_gemma":0.011727501,"teacher_disagreement_score":0.0063129645,"about_ca_system_score_codex":0.0007399251,"about_ca_system_score_gemma":0.001708936,"threshold_uncertainty_score":0.01255244},"labels":[],"label_agreement":null},{"id":"W3177727271","doi":"10.1109/access.2021.3135807","title":"Comparative Analysis of Word Embeddings in Assessing Semantic Similarity of Complex Sentences","year":2021,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Lakehead University","keywords":"Computer science; Natural language processing; Artificial intelligence; Semantic similarity; Word embedding; Sentence; Benchmark (surveying); Readability; Word (group theory); Similarity (geometry); Transformer; Embedding; Linguistics","score_opus":0.14327590750997854,"score_gpt":0.4075290179436385,"score_spread":0.26425311043365995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3177727271","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92346555,0.0037857452,0.060006388,0.00039692194,0.00031126456,0.00018702698,0.007245046,0.0017039081,0.0028981632],"genre_scores_gemma":[0.93798375,0.0008072317,0.03885845,0.00007564388,0.0001332785,0.00013762453,0.021182906,0.00014280094,0.00067830604],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957841,0.0021036353,0.00050774176,0.0007941853,0.00065953674,0.00015085784],"domain_scores_gemma":[0.9826909,0.01154502,0.0013860869,0.0018829993,0.0021040658,0.00039095353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048167165,0.0012023087,0.0007706962,0.0044632414,0.00041893285,0.0014866161,0.00070575054,0.00085734215,0.00092272816],"category_scores_gemma":[0.026009249,0.00016936488,0.0008079581,0.0031972087,0.0005393604,0.0032257913,0.001389438,0.000822308,0.0007220388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039439113,0.0017042173,0.16348122,0.0037007295,0.00250628,0.0009163981,0.0026660399,0.09962275,0.045241423,0.006964801,0.028355462,0.6408968],"study_design_scores_gemma":[0.00016023312,0.0029539221,0.10534372,0.00024226042,0.0006810729,0.0015774164,0.0030871576,0.8294414,0.03137305,0.013616016,0.011303791,0.0002199697],"about_ca_topic_score_codex":0.0013967388,"about_ca_topic_score_gemma":0.002819163,"teacher_disagreement_score":0.0048167165,"about_ca_system_score_codex":0.00047934774,"about_ca_system_score_gemma":0.00048265437,"threshold_uncertainty_score":0.025473535},"labels":[],"label_agreement":null},{"id":"W3178029650","doi":"10.1016/j.jbi.2021.103864","title":"Development of a generalizable natural language processing pipeline to extract physician-reported pain from clinical reports: Generated using publicly-available datasets and tested on institutional clinical reports for cancer patients with bone metastases","year":2021,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; McGill University Health Centre","funders":"","keywords":"Pipeline (software); Medicine; Computer science","score_opus":0.07284276618582329,"score_gpt":0.3608962595140042,"score_spread":0.2880534933281809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3178029650","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1124898,0.0016683132,0.3911177,0.00254626,0.0006758941,0.002959282,0.28329363,0.2017421,0.0035070563],"genre_scores_gemma":[0.13294154,0.00063053204,0.43795577,0.000515666,0.0001882087,0.0016534778,0.42123836,0.0018159981,0.0030604806],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988939,0.0001746908,0.00017026029,0.0004423078,0.00021063857,0.00010812736],"domain_scores_gemma":[0.9971245,0.0014399681,0.000214034,0.00032205222,0.00075736176,0.00014195232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018201545,0.0021710652,0.0008432942,0.003983861,0.0006290567,0.0015015639,0.0013859438,0.0011408459,0.004800575],"category_scores_gemma":[0.0061622513,0.00072407257,0.002269258,0.002128907,0.00028096483,0.0013912845,0.0015879619,0.0015251464,0.005852227],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013997906,0.001528838,0.04602175,0.0034635684,0.001271561,0.0025095777,0.0014706043,0.024547512,0.06878597,0.0026355057,0.28872663,0.55763876],"study_design_scores_gemma":[0.00069244334,0.0007226102,0.07194856,0.0004208128,0.00084926136,0.0027400746,0.0017505231,0.70687854,0.07021215,0.012227784,0.13123919,0.00031805996],"about_ca_topic_score_codex":0.018085452,"about_ca_topic_score_gemma":0.025574962,"teacher_disagreement_score":0.018085452,"about_ca_system_score_codex":0.000806859,"about_ca_system_score_gemma":0.0032283757,"threshold_uncertainty_score":0.035960376},"labels":[],"label_agreement":null},{"id":"W3179689435","doi":"","title":"A Joint Model for Question Answering and Question Generation","year":2017,"lang":"en","type":"article","venue":"International Conference on Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Novelty; Question answering; Artificial intelligence; Generative model; Comprehension; Perspective (graphical); Generative grammar; Joint (building); Sequence (biology); Natural language processing; Ask price; Sequence labeling; Machine learning; Programming language; Task (project management); Engineering","score_opus":0.11396602430690275,"score_gpt":0.3511012088271431,"score_spread":0.23713518452024032,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3179689435","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015905457,0.00040065328,0.97673756,0.0013288601,0.00007715186,0.00013903728,0.0003860239,0.0017865837,0.0032386691],"genre_scores_gemma":[0.6286121,0.0006265665,0.35232916,0.0008277804,0.00031737107,0.0007940082,0.0019945502,0.000499165,0.013999326],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983296,0.0007720788,0.000080581456,0.00048410238,0.00020645541,0.00012721618],"domain_scores_gemma":[0.99474156,0.0038111121,0.0002432752,0.0005298797,0.0004946989,0.00017937424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030777904,0.0009631861,0.001082689,0.0016041175,0.00068072916,0.0022332177,0.0029320645,0.0027597058,0.0055998606],"category_scores_gemma":[0.010638429,0.0009407821,0.0021655457,0.0012438992,0.0015378351,0.004706353,0.0020392733,0.0030700427,0.0023631689],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056483573,0.00044038708,0.007075779,0.0005062467,0.00034253596,0.00063319446,0.0029384575,0.4236778,0.014166462,0.29665768,0.013262873,0.2397337],"study_design_scores_gemma":[0.00002591389,0.000046518,0.00027039283,0.000014250418,0.000040807445,0.000120911805,0.00003526575,0.92421275,0.001150052,0.07180949,0.0022525536,0.000021014166],"about_ca_topic_score_codex":0.004984634,"about_ca_topic_score_gemma":0.0053216317,"teacher_disagreement_score":0.0055998606,"about_ca_system_score_codex":0.0013447918,"about_ca_system_score_gemma":0.0018656009,"threshold_uncertainty_score":0.018733382},"labels":[],"label_agreement":null},{"id":"W3180273654","doi":"10.2139/ssrn.3708327","title":"Trends in COVID-19 Publications: Streamlining Research Using NLP and LDA","year":2020,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trinity College","funders":"","keywords":"Coronavirus disease 2019 (COVID-19); Natural language processing; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Artificial intelligence; 2019-20 coronavirus outbreak; Computer science; Medicine; Virology","score_opus":0.16578031830386628,"score_gpt":0.4051460986215873,"score_spread":0.23936578031772102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3180273654","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45399255,0.12502173,0.049534436,0.049331088,0.0065465695,0.00092978484,0.20849301,0.0075657438,0.09858512],"genre_scores_gemma":[0.6894463,0.061875153,0.0647916,0.0030723456,0.006333028,0.0009484849,0.1539824,0.002370436,0.01718029],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98844624,0.00216016,0.0024144242,0.0020723066,0.0042891796,0.0006177535],"domain_scores_gemma":[0.88588464,0.057987887,0.013760369,0.0059330366,0.03199306,0.004441014],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.012847404,0.00069737545,0.0011758578,0.061024506,0.0013908102,0.012863669,0.0013750067,0.0013378226,0.006560401],"category_scores_gemma":[0.074792095,0.00046407126,0.0013712138,0.07684485,0.0009209566,0.009276573,0.0028510806,0.0021359539,0.006651492],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005276669,0.00021190212,0.24845432,0.0058015827,0.0004396603,0.0003528563,0.0052178004,0.0014445205,0.0050384933,0.015297026,0.10403737,0.61317676],"study_design_scores_gemma":[0.00011976426,0.00039830402,0.4218209,0.0037701859,0.000946842,0.0012697799,0.0115521215,0.017658569,0.008869709,0.020051995,0.5133393,0.00020247241],"about_ca_topic_score_codex":0.008364301,"about_ca_topic_score_gemma":0.012053854,"teacher_disagreement_score":0.9871526,"about_ca_system_score_codex":0.0020978735,"about_ca_system_score_gemma":0.0061298423,"threshold_uncertainty_score":0.06794441},"labels":[],"label_agreement":null},{"id":"W3183195595","doi":"10.1162/tacl_a_00386","title":"Context-aware Adversarial Training for Name Regularity Bias in Named Entity Recognition","year":2021,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Huawei Technologies (Canada)","funders":"","keywords":"Adversarial system; Focus (optics); Testbed; Named-entity recognition; Training set; Noise (video); Training (meteorology)","score_opus":0.0861111343750003,"score_gpt":0.2954547619652698,"score_spread":0.20934362759026953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3183195595","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27179238,0.0017042925,0.71447563,0.00088328955,0.0003390352,0.000092918504,0.0005567345,0.0060043093,0.004151265],"genre_scores_gemma":[0.93791556,0.00016918645,0.059003476,0.00030019932,0.00006552102,0.0000501406,0.0006612039,0.00024725156,0.0015874268],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99902916,0.00047615482,0.00004290468,0.0002603734,0.00009831541,0.00009307304],"domain_scores_gemma":[0.99483985,0.003737454,0.00021779668,0.0008696502,0.00023090096,0.00010440735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029144764,0.0009759403,0.00081585214,0.0003711851,0.00036974327,0.0005509856,0.0016012815,0.0011470203,0.0017256653],"category_scores_gemma":[0.00915727,0.00048390325,0.0004993653,0.00034321786,0.0010111285,0.0018332405,0.0017110264,0.0028071604,0.00065641035],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037496476,0.00014750162,0.0025842236,0.00012650179,0.00010476592,0.00013006994,0.00010567938,0.8866991,0.011521505,0.003537005,0.0033603932,0.09130825],"study_design_scores_gemma":[0.000006172966,0.000038823982,0.00023046353,0.000010525906,0.000008682974,0.000022450064,0.000008732326,0.99349606,0.0040181824,0.0018034076,0.00034913188,0.0000074151726],"about_ca_topic_score_codex":0.002651697,"about_ca_topic_score_gemma":0.0036885624,"teacher_disagreement_score":0.0029144764,"about_ca_system_score_codex":0.00054259936,"about_ca_system_score_gemma":0.0005648591,"threshold_uncertainty_score":0.0154134035},"labels":[],"label_agreement":null},{"id":"W3183335868","doi":"10.1007/978-3-030-79457-6_48","title":"Collapsed Gibbs Sampling of Beta-Liouville Multinomial for Short Text Clustering","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Gibbs sampling; Multinomial distribution; Cluster analysis; Sampling (signal processing); Social media; Data mining; Artificial intelligence; World Wide Web; Statistics; Mathematics; Bayesian probability","score_opus":0.04681323560420386,"score_gpt":0.2805330171014423,"score_spread":0.23371978149723843,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3183335868","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012834202,0.0005300724,0.98443896,0.00024369276,0.000100242585,0.00012498435,0.00037319906,0.0006283327,0.0007263072],"genre_scores_gemma":[0.32878986,0.0011747669,0.6390702,0.000709511,0.000870673,0.0017187443,0.009291771,0.0016231012,0.016751416],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99430907,0.0034648916,0.0002711318,0.0009961582,0.0006086827,0.00035009722],"domain_scores_gemma":[0.9764355,0.018714568,0.0004538731,0.0022196274,0.001498102,0.0006784213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0081256,0.001130058,0.0030129978,0.0041106753,0.0024434985,0.003685984,0.005766449,0.0036238388,0.007979943],"category_scores_gemma":[0.035409037,0.001494093,0.0029784646,0.005043009,0.0025175489,0.0045042303,0.0041498346,0.003395226,0.002950497],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090922887,0.00037941284,0.0043049953,0.0007000974,0.0004529217,0.00045008867,0.0015908201,0.42009377,0.0059076305,0.33141333,0.01938857,0.2144091],"study_design_scores_gemma":[0.000024186173,0.000017286065,0.00024285409,0.000023755902,0.000016034546,0.000039480714,0.000039411836,0.93105847,0.00038360307,0.06685813,0.0012771018,0.00001961643],"about_ca_topic_score_codex":0.010174997,"about_ca_topic_score_gemma":0.017482597,"teacher_disagreement_score":0.010174997,"about_ca_system_score_codex":0.002663674,"about_ca_system_score_gemma":0.0023991957,"threshold_uncertainty_score":0.042972803},"labels":[],"label_agreement":null},{"id":"W3183902696","doi":"10.5539/cis.v14n3p78","title":"On Bi-gram Graph Attributes","year":2021,"lang":"en","type":"article","venue":"Computer and Information Science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Scalability; Graph; n-gram; Theoretical computer science; Representation (politics); Information retrieval; Natural language processing; Language model; Database","score_opus":0.023152165114462477,"score_gpt":0.24848098079430916,"score_spread":0.2253288156798467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3183902696","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012631505,0.000502324,0.9765089,0.0009817518,0.00014894553,0.00015797197,0.002000235,0.0018554328,0.0052128267],"genre_scores_gemma":[0.20928964,0.00093922624,0.7756658,0.0005559445,0.00046515756,0.0005683289,0.006295869,0.0011450064,0.005074986],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99574125,0.0015609991,0.00023263796,0.0012526914,0.0009957529,0.00021668118],"domain_scores_gemma":[0.988822,0.0052064504,0.0009250198,0.0030324748,0.001622373,0.00039183174],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002701161,0.0011346185,0.0012299306,0.011304807,0.0021202858,0.0057228752,0.0020041838,0.0019205824,0.0056206067],"category_scores_gemma":[0.021087104,0.0006032323,0.0015826245,0.013437747,0.003220174,0.012801047,0.0043747677,0.0031576124,0.002960659],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002253621,0.00016377878,0.004280639,0.0005286985,0.0001375567,0.00035791937,0.00218293,0.020101441,0.0074035213,0.7112738,0.012897257,0.24044709],"study_design_scores_gemma":[0.000010818217,0.000046208792,0.0014989473,0.00010454467,0.000047189616,0.00028261245,0.0007844621,0.14457208,0.0025033634,0.8144077,0.035678584,0.00006346091],"about_ca_topic_score_codex":0.0047278013,"about_ca_topic_score_gemma":0.005430365,"teacher_disagreement_score":0.011304807,"about_ca_system_score_codex":0.0016883919,"about_ca_system_score_gemma":0.0018480485,"threshold_uncertainty_score":0.018802762},"labels":[],"label_agreement":null},{"id":"W3184197775","doi":"10.1145/3446390","title":"Chinese Emotional Dialogue Response Generation via Reinforcement Learning","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Internet Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Natural Science Foundation of China","keywords":"Computer science; Reinforcement learning; Expression (computer science); Artificial intelligence; Function (biology); Quality (philosophy); Key (lock); Process (computing); Reinforcement; Machine learning; Psychology; Social psychology","score_opus":0.018306925265429815,"score_gpt":0.26179627202930583,"score_spread":0.243489346763876,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3184197775","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11732491,0.00023779215,0.8753578,0.00026774168,0.00008472396,0.0002359459,0.000053145857,0.0023071235,0.004130812],"genre_scores_gemma":[0.891035,0.000079626036,0.10505367,0.00016042344,0.000027802014,0.00027506755,0.000107918706,0.000092089,0.003168287],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99908674,0.00040741332,0.00004201338,0.00024476863,0.00012773808,0.00009119065],"domain_scores_gemma":[0.99888176,0.00065695005,0.00007038366,0.00006852692,0.0002596291,0.00006275486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014684923,0.0007942754,0.0006651655,0.00033945774,0.00036963914,0.00043716916,0.00087371434,0.00062332087,0.0020208538],"category_scores_gemma":[0.0036286057,0.00021924525,0.00043713272,0.00023831187,0.00049282244,0.0006371985,0.0007341355,0.0006436804,0.00038713918],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008082118,0.00045185458,0.0035560136,0.00031614825,0.00009699278,0.00042868787,0.001024463,0.44599673,0.035579048,0.009401543,0.0045371526,0.49780318],"study_design_scores_gemma":[0.000036379304,0.0000680064,0.00024526528,0.0000035090368,0.00001297202,0.000035022822,0.00002535679,0.99449,0.0032837796,0.0012698326,0.00052050553,0.000009357199],"about_ca_topic_score_codex":0.0025716727,"about_ca_topic_score_gemma":0.0017382527,"teacher_disagreement_score":0.0025716727,"about_ca_system_score_codex":0.00050104316,"about_ca_system_score_gemma":0.00059170695,"threshold_uncertainty_score":0.007766247},"labels":[],"label_agreement":null},{"id":"W3184501572","doi":"10.18653/v1/2021.starsem-1.29","title":"Evaluating a Joint Training Approach for Learning Cross-lingual Embeddings with Sub-word Information without Parallel Corpora on Lower-resource Languages","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Natural language processing; Word (group theory); Artificial intelligence; Lexicon; Joint (building); Parallel corpora; Resource (disambiguation); Vocabulary; Extension (predicate logic); Machine translation; Linguistics; Programming language","score_opus":0.08097096599880486,"score_gpt":0.34043831840685834,"score_spread":0.25946735240805346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3184501572","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39756364,0.0051655336,0.5595452,0.0009165947,0.0008857753,0.0006454279,0.0027414772,0.016789109,0.015747124],"genre_scores_gemma":[0.654181,0.0010217074,0.31421947,0.000571733,0.00021469165,0.00078099995,0.017656283,0.0017418041,0.009612257],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99443555,0.0026475394,0.000427516,0.0015311717,0.0006333678,0.00032483466],"domain_scores_gemma":[0.98859566,0.0061699636,0.00032275615,0.0022865797,0.0022682035,0.0003567417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00955563,0.0027755573,0.0013015957,0.0019753003,0.00097541366,0.002000061,0.0020094344,0.0024002404,0.004104203],"category_scores_gemma":[0.018227156,0.0009781492,0.0014220814,0.0025497298,0.00096288155,0.0076234606,0.0043888763,0.0030070234,0.00276998],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014587938,0.0014449271,0.013025017,0.0007536417,0.0011028136,0.0002677888,0.00068829657,0.18505102,0.020191738,0.0037276,0.019785887,0.7525025],"study_design_scores_gemma":[0.00021898131,0.0008658649,0.0045542805,0.00008826584,0.00027600667,0.0002563946,0.0004995508,0.96085465,0.01911871,0.0046813306,0.008502169,0.000083917745],"about_ca_topic_score_codex":0.009814663,"about_ca_topic_score_gemma":0.016339676,"teacher_disagreement_score":0.009814663,"about_ca_system_score_codex":0.0010630175,"about_ca_system_score_gemma":0.0018169545,"threshold_uncertainty_score":0.05053556},"labels":[],"label_agreement":null},{"id":"W3184918446","doi":"10.18653/v1/2021.repl4nlp-1.17","title":"In-Batch Negatives for Knowledge Distillation with Tightly-Coupled Teachers for Dense Retrieval","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":119,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Encoder; Computer science; Distillation; Ranking (information retrieval); Artificial intelligence; Knowledge transfer; Learning to rank; Information retrieval; Machine learning; Knowledge management","score_opus":0.02914580604038951,"score_gpt":0.28625465068007594,"score_spread":0.2571088446396864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3184918446","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014831798,0.00031779835,0.9740174,0.00040832377,0.00010026603,0.00008857016,0.00031737555,0.007362775,0.0025557622],"genre_scores_gemma":[0.4641142,0.0002783151,0.5122764,0.00087721506,0.00022581554,0.00045766743,0.0023109424,0.0018002143,0.017659187],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99919015,0.00025466588,0.000040690815,0.00023546006,0.00017450262,0.00010454229],"domain_scores_gemma":[0.99830014,0.00081502146,0.000114225564,0.00044715617,0.00022083058,0.00010260188],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014559196,0.0017552694,0.0011728599,0.0005612797,0.0007700715,0.0014795504,0.0031573602,0.0016408624,0.007955279],"category_scores_gemma":[0.0066994485,0.00086597895,0.00080938137,0.0006494004,0.0013249863,0.0039592697,0.003112664,0.0038018618,0.004338209],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007830102,0.0004461779,0.0016332865,0.00045508213,0.00012971583,0.00030542337,0.00045107966,0.32855093,0.023307327,0.047659326,0.026030805,0.5702479],"study_design_scores_gemma":[0.000037791488,0.00007995897,0.00010720866,0.000018345401,0.000019747398,0.000060822098,0.000027501655,0.97097653,0.007821436,0.017364286,0.0034672988,0.000019131805],"about_ca_topic_score_codex":0.0054324074,"about_ca_topic_score_gemma":0.014575297,"teacher_disagreement_score":0.007955279,"about_ca_system_score_codex":0.0011821976,"about_ca_system_score_gemma":0.0018250055,"threshold_uncertainty_score":0.026613057},"labels":[],"label_agreement":null},{"id":"W3184956571","doi":"10.18653/v1/2021.sigdial-1.50","title":"Do Encoder Representations of Generative Dialogue Models have sufficient summary of the Information about the task ?","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Encoder; Transformer; Language model; Artificial intelligence; Representation (politics); Utterance; Generative grammar; Natural language processing; Task (project management); Generative model; Machine learning","score_opus":0.03447748161833202,"score_gpt":0.2667199454485412,"score_spread":0.2322424638302092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3184956571","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25254956,0.0012978677,0.7280654,0.0018340313,0.00023103754,0.0002458887,0.0011422979,0.004657345,0.00997672],"genre_scores_gemma":[0.9294724,0.0003145075,0.06565569,0.0003147184,0.00006127416,0.00020835247,0.0013446372,0.0004310279,0.0021974551],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99677974,0.0020837102,0.0001067259,0.0005239621,0.00029082855,0.00021504027],"domain_scores_gemma":[0.9811334,0.013510682,0.0007869212,0.0026011071,0.0013998513,0.00056799804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005325685,0.001082193,0.0009824941,0.00054552354,0.00036040274,0.0022761954,0.0011288099,0.0016478475,0.0032916917],"category_scores_gemma":[0.04044022,0.000680413,0.00060332543,0.00041052734,0.0007717295,0.005456813,0.0011815515,0.0022534227,0.0020236378],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026496323,0.000689818,0.019809555,0.0015982551,0.0005795575,0.00035915844,0.002636954,0.3759925,0.042391557,0.037299655,0.009279564,0.50671375],"study_design_scores_gemma":[0.000062584644,0.00025900122,0.0023751156,0.00010353931,0.00007275597,0.00010483485,0.00030907147,0.9654885,0.007536755,0.021763809,0.0018911593,0.000032966396],"about_ca_topic_score_codex":0.0029631145,"about_ca_topic_score_gemma":0.0045737512,"teacher_disagreement_score":0.005325685,"about_ca_system_score_codex":0.0009570988,"about_ca_system_score_gemma":0.001315496,"threshold_uncertainty_score":0.02816528},"labels":[],"label_agreement":null},{"id":"W3185670382","doi":"10.1145/3462757.3466147","title":"CriminelBART","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université Laval","funders":"","keywords":"Computer science; Transformer; Vocabulary; Artificial intelligence; Natural language processing; Architecture; Language model; Natural language; Comprehension; Natural language understanding; Linguistics; Programming language","score_opus":0.04965968158271803,"score_gpt":0.2612948729727764,"score_spread":0.2116351913900584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3185670382","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.069268,0.006977573,0.36335388,0.011119723,0.0042622145,0.0011387181,0.03156654,0.14881156,0.36350182],"genre_scores_gemma":[0.291318,0.0034754085,0.2926227,0.0036248348,0.0011525629,0.00074718986,0.06443477,0.020804146,0.32182047],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99884486,0.00011699079,0.000035513764,0.00040425538,0.00046088206,0.0001374444],"domain_scores_gemma":[0.9987375,0.00034759202,0.000053110518,0.00034355163,0.00036999956,0.000148202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010091033,0.0014339621,0.0005137561,0.0018345148,0.0013745218,0.0019862328,0.0014889835,0.0011302975,0.056265682],"category_scores_gemma":[0.004736861,0.000521192,0.00055990566,0.0015144756,0.0006639482,0.0029494774,0.001977092,0.0018359772,0.032425627],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033049635,0.00014368277,0.0030196926,0.00022665421,0.00005027848,0.0005039949,0.00033094126,0.0063743345,0.0070676957,0.013718949,0.3129876,0.65524566],"study_design_scores_gemma":[0.00009245244,0.00012594183,0.0034996711,0.0001034021,0.0000532633,0.00085836527,0.00013619158,0.071537,0.018565074,0.016215615,0.88871497,0.00009790903],"about_ca_topic_score_codex":0.029834755,"about_ca_topic_score_gemma":0.040360007,"teacher_disagreement_score":0.056265682,"about_ca_system_score_codex":0.002097124,"about_ca_system_score_gemma":0.00275492,"threshold_uncertainty_score":0.18822747},"labels":[],"label_agreement":null},{"id":"W3185919331","doi":"10.22215/etd/2021-14497","title":"Tales of a Coronavirus Pandemic: Topic Modelling with Short-Text Data","year":2021,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Latent Dirichlet allocation; Topic model; Computer science; Term (time); Coronavirus disease 2019 (COVID-19); Natural language processing; Information retrieval; Pandemic; Artificial intelligence; Data science; Medicine","score_opus":0.15631076015034184,"score_gpt":0.3371977118474921,"score_spread":0.18088695169715024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3185919331","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5790862,0.011345272,0.3814184,0.011301974,0.00052556075,0.00041767725,0.005490657,0.0011680509,0.009246206],"genre_scores_gemma":[0.87441146,0.0045627877,0.10859922,0.00039913174,0.00065799366,0.00024916927,0.005854692,0.00014836041,0.0051171337],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99877185,0.000703593,0.00006336747,0.00020533863,0.00019459211,0.00006134446],"domain_scores_gemma":[0.9821217,0.015948279,0.0005623182,0.0005274934,0.0006114994,0.00022872955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048908126,0.0005960958,0.0005360605,0.002560201,0.0007076939,0.0025893878,0.00079833914,0.0012446032,0.0016959542],"category_scores_gemma":[0.01954673,0.00044025382,0.0010997558,0.0024808685,0.0005775647,0.003263726,0.00090430677,0.00240022,0.0008544311],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010353628,0.0006140769,0.06828276,0.0012020756,0.0008111995,0.000932947,0.008822118,0.36932886,0.00526064,0.028393114,0.026250148,0.4890667],"study_design_scores_gemma":[0.000029980194,0.00014497565,0.020826492,0.00014610017,0.000112162874,0.0001888273,0.0018520093,0.92002696,0.0017342811,0.041967824,0.01288667,0.00008379464],"about_ca_topic_score_codex":0.005607553,"about_ca_topic_score_gemma":0.005666941,"teacher_disagreement_score":0.005607553,"about_ca_system_score_codex":0.00084285275,"about_ca_system_score_gemma":0.00062298204,"threshold_uncertainty_score":0.025865436},"labels":[],"label_agreement":null},{"id":"W3185922622","doi":"10.1145/3462757.3466102","title":"Using transformers to improve answer retrieval for legal questions","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Automatic summarization; Question answering; Transformer; Support vector machine; Artificial intelligence; Machine translation; Machine learning; Mean reciprocal rank; Natural language processing; Information retrieval; Engineering; Voltage","score_opus":0.041296470959886984,"score_gpt":0.3094840678173827,"score_spread":0.2681875968574957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3185922622","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24470702,0.0049746865,0.65181893,0.0023750504,0.0006948955,0.00067573966,0.0033882787,0.07396164,0.017403755],"genre_scores_gemma":[0.76436746,0.0012127206,0.2111887,0.0007916379,0.00037887399,0.00019952543,0.011358175,0.0008500522,0.009652843],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984273,0.00046488395,0.00013520704,0.00040126633,0.00041447754,0.00015693536],"domain_scores_gemma":[0.99677104,0.0015793617,0.00019310528,0.0005265706,0.0007974452,0.00013245075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029526227,0.0012369656,0.0009123451,0.0035829267,0.00060493016,0.0018931563,0.0012446939,0.0010706598,0.0051308996],"category_scores_gemma":[0.012005687,0.00029674367,0.00091502204,0.0018945164,0.0005131855,0.005900378,0.0016178732,0.0014286512,0.0051357723],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079481036,0.0005750402,0.008975206,0.00047569865,0.00012480504,0.00018767394,0.0005509446,0.021655617,0.036647476,0.014437792,0.044889096,0.8706859],"study_design_scores_gemma":[0.00016621882,0.0006743654,0.0036527307,0.000049015547,0.00021105152,0.0004066449,0.00040243586,0.8884844,0.04952611,0.024813697,0.031550083,0.00006325021],"about_ca_topic_score_codex":0.0059175505,"about_ca_topic_score_gemma":0.009533649,"teacher_disagreement_score":0.0059175505,"about_ca_system_score_codex":0.0010270874,"about_ca_system_score_gemma":0.0016422425,"threshold_uncertainty_score":0.017164528},"labels":[],"label_agreement":null},{"id":"W3186310017","doi":"10.18653/v1/2021.semeval-1.48","title":"CLaC-np at SemEval-2021 Task 8: Dependency DGCNN","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; SemEval; Dependency (UML); Artificial intelligence; Natural language processing; Lexical analysis; Preprocessor; Task (project management); Variety (cybernetics); ENCODE; Graph; Dependency graph; Theoretical computer science","score_opus":0.019222937934521723,"score_gpt":0.24289961460019102,"score_spread":0.2236766766656693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3186310017","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34469318,0.005241951,0.16322936,0.007574435,0.006876114,0.0028227863,0.236892,0.11879846,0.1138717],"genre_scores_gemma":[0.4871244,0.0004514787,0.12142406,0.0025587867,0.00056007114,0.002244871,0.33140117,0.0060191764,0.0482159],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982002,0.00039189006,0.00008196472,0.0008064816,0.00023846122,0.00028092752],"domain_scores_gemma":[0.9965275,0.0015891881,0.00011904927,0.0008209575,0.000715475,0.00022777302],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022763414,0.0036442932,0.0014667021,0.001077169,0.001421472,0.0021825181,0.0029916223,0.0049140975,0.035302874],"category_scores_gemma":[0.009892116,0.00069791806,0.0017636468,0.0011014558,0.0006801407,0.004484251,0.002623425,0.005133244,0.021704383],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015065024,0.00078640686,0.005059825,0.0012758505,0.00031056078,0.0014475284,0.00038483372,0.02201171,0.013159236,0.00825463,0.6602568,0.2855462],"study_design_scores_gemma":[0.0009357285,0.0008850922,0.011347876,0.00042329068,0.00024415302,0.0018241723,0.0008245463,0.6474656,0.052279644,0.04612514,0.23741868,0.00022615549],"about_ca_topic_score_codex":0.014839328,"about_ca_topic_score_gemma":0.020006845,"teacher_disagreement_score":0.035302874,"about_ca_system_score_codex":0.0017770157,"about_ca_system_score_gemma":0.001994205,"threshold_uncertainty_score":0.11809987},"labels":[],"label_agreement":null},{"id":"W3187346832","doi":"10.24963/ijcai.2021/544","title":"Hierarchical Modeling of Label Dependency and Label Noise in Fine-grained Entity Typing","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China","keywords":"Computer science; Dependency (UML); Tree (set theory); Hierarchy; Benchmark (surveying); Artificial intelligence; Noise (video); Sentence; Natural language processing; Confusion; Machine learning","score_opus":0.03658541278609571,"score_gpt":0.2701452910301649,"score_spread":0.23355987824406918,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3187346832","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09589938,0.00055285473,0.9004149,0.0003168443,0.000050565362,0.00006881711,0.0005129995,0.0012494507,0.0009342162],"genre_scores_gemma":[0.80941117,0.00043871292,0.18167299,0.00027250763,0.00014998643,0.00020337066,0.0027403834,0.0004654584,0.0046454133],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99811816,0.00061904266,0.00010944444,0.00068470393,0.0002545113,0.00021414923],"domain_scores_gemma":[0.9888602,0.0076990323,0.00095147634,0.001226102,0.00097350916,0.00028979738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039125504,0.0011297329,0.0013361108,0.0021100885,0.0012542115,0.0014865707,0.0025572353,0.0018172667,0.0015015483],"category_scores_gemma":[0.011832825,0.0008630242,0.0014705565,0.0025491016,0.001195663,0.005423707,0.0020915351,0.00323303,0.0008408862],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007019391,0.0005020989,0.03879876,0.00039574102,0.00025066058,0.00075457396,0.0035512634,0.56860524,0.017160125,0.063489504,0.008756305,0.29703373],"study_design_scores_gemma":[0.0000075755897,0.00002158272,0.0016822414,0.000011932292,0.000024901688,0.00005335466,0.000043653254,0.97596675,0.0012618394,0.020171402,0.0007404189,0.0000143457855],"about_ca_topic_score_codex":0.014293134,"about_ca_topic_score_gemma":0.02924433,"teacher_disagreement_score":0.014293134,"about_ca_system_score_codex":0.0014202333,"about_ca_system_score_gemma":0.0016798138,"threshold_uncertainty_score":0.028419852},"labels":[],"label_agreement":null},{"id":"W3188561342","doi":"10.5194/egusphere-egu22-13024","title":"Hydrology Research Articles Are Becoming More Interdisciplinary","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Summer camp; Adenylate kinase; Chemistry; Biochemistry; Sociology; Enzyme","score_opus":0.16750673123075643,"score_gpt":0.41821566907322916,"score_spread":0.2507089378424727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3188561342","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8020435,0.06402962,0.0053031067,0.013386277,0.0013654138,0.0004856534,0.028170384,0.0007736422,0.084442385],"genre_scores_gemma":[0.91133547,0.03037765,0.011493286,0.0032443479,0.0032669695,0.0006730883,0.025806956,0.00033925535,0.0134629095],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9788533,0.003956131,0.004725547,0.0033847396,0.008332627,0.0007477525],"domain_scores_gemma":[0.8494747,0.06041817,0.05727748,0.0049652657,0.021598296,0.006266088],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.014963644,0.0005254223,0.0011123599,0.04082887,0.0014500681,0.008399086,0.0010988336,0.00095538923,0.012153689],"category_scores_gemma":[0.051574737,0.0005045942,0.0012676712,0.060160052,0.0012445748,0.006550048,0.0051643173,0.0010313075,0.003374224],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065929216,0.00018133041,0.4758249,0.015832027,0.0016876026,0.0022047968,0.028790927,0.001363954,0.025550634,0.008970886,0.045627307,0.39330631],"study_design_scores_gemma":[0.000060672075,0.00013759467,0.8115502,0.00095384195,0.00045858065,0.001163476,0.011732329,0.00059785333,0.0022365886,0.0048164716,0.16621,0.000082390725],"about_ca_topic_score_codex":0.0024907805,"about_ca_topic_score_gemma":0.0044668936,"teacher_disagreement_score":0.9850364,"about_ca_system_score_codex":0.0020787884,"about_ca_system_score_gemma":0.0019843595,"threshold_uncertainty_score":0.07913625},"labels":[],"label_agreement":null},{"id":"W3189064355","doi":"10.1007/s40747-021-00482-y","title":"Modeling multi-prototype Chinese word representation learning for word similarity","year":2021,"lang":"en","type":"article","venue":"Complex & Intelligent Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Word (group theory); Natural language processing; Artificial intelligence; Similarity (geometry); Polysemy; Semantic similarity; Word embedding; Representation (politics); Task (project management); Embedding; Mathematics","score_opus":0.160357679159926,"score_gpt":0.35882770318270113,"score_spread":0.19847002402277514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3189064355","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21749154,0.0011148478,0.7769193,0.00043292734,0.0001103572,0.00015205378,0.00050085667,0.0014524253,0.0018256793],"genre_scores_gemma":[0.8891519,0.00037188342,0.10611277,0.00012989425,0.00007581521,0.0001923638,0.0013054081,0.000098394215,0.0025616814],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987087,0.00034938962,0.000091507434,0.00055654295,0.00017836413,0.00011539791],"domain_scores_gemma":[0.99852055,0.0007311245,0.00017894674,0.00019438607,0.0003212332,0.000053725147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015852548,0.00089387345,0.0010545513,0.0024309277,0.00042871185,0.0010191164,0.0018111522,0.0010268624,0.0015249723],"category_scores_gemma":[0.004955462,0.00035192896,0.0011770211,0.0023940997,0.000529425,0.0031330495,0.0011561932,0.0011611648,0.0005338475],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003694438,0.00042455393,0.01150612,0.00032210065,0.0003392215,0.00028890377,0.0004421151,0.30775,0.009703805,0.01787643,0.006869374,0.644108],"study_design_scores_gemma":[0.000008388449,0.0000333958,0.0005592246,0.0000055477317,0.000022812674,0.000030416475,0.000019889883,0.99505794,0.00074098574,0.0032035932,0.00030929619,0.000008542524],"about_ca_topic_score_codex":0.0071277907,"about_ca_topic_score_gemma":0.0059804805,"teacher_disagreement_score":0.0071277907,"about_ca_system_score_codex":0.0010158094,"about_ca_system_score_gemma":0.00092036184,"threshold_uncertainty_score":0.014172614},"labels":[],"label_agreement":null},{"id":"W3190588453","doi":"10.3390/data6080084","title":"The Automatic Detection of Dataset Names in Scientific Articles","year":2021,"lang":"en","type":"article","venue":"Data","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Canadian Institute of Steel Construction","keywords":"Computer science; Named-entity recognition; Annotation; Task (project management); Natural language processing; Sentence; Feature (linguistics); Set (abstract data type); Artificial intelligence; Code (set theory); Information retrieval; Baseline (sea); Linguistics; Programming language","score_opus":0.06923547144482206,"score_gpt":0.2931762461491873,"score_spread":0.22394077470436524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3190588453","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53058195,0.012373835,0.21335404,0.0028768415,0.0023467157,0.0010149715,0.18580905,0.028967366,0.022675203],"genre_scores_gemma":[0.41023466,0.0022471435,0.26612952,0.0005339732,0.0008666245,0.00071203633,0.3096562,0.0016038368,0.008016094],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99537885,0.0011670991,0.0006589513,0.0015442661,0.0010046565,0.00024617446],"domain_scores_gemma":[0.9716234,0.013969642,0.004356189,0.0028594781,0.0063418155,0.0008493278],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.004893071,0.0009207206,0.00063640863,0.009043719,0.0011093831,0.0020217877,0.0011502025,0.0012828917,0.0017361391],"category_scores_gemma":[0.018701801,0.00035127028,0.00081843993,0.006277673,0.0006306232,0.0036914009,0.0015812117,0.0011109636,0.0043962854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000992422,0.0005106914,0.13300802,0.00788343,0.00037154602,0.0020208063,0.0040747975,0.009351143,0.10167912,0.010511164,0.21457756,0.51501924],"study_design_scores_gemma":[0.00010685937,0.00072204985,0.20755507,0.0007488759,0.0004840919,0.004381821,0.0035126815,0.1442869,0.19515064,0.01594427,0.4267844,0.00032237082],"about_ca_topic_score_codex":0.002664386,"about_ca_topic_score_gemma":0.0067647444,"teacher_disagreement_score":0.99510694,"about_ca_system_score_codex":0.0008263946,"about_ca_system_score_gemma":0.001561098,"threshold_uncertainty_score":0.025877297},"labels":[],"label_agreement":null},{"id":"W3192525318","doi":"10.20381/ruor-23341","title":"Towards the Automatic Classification of Student Answers to Open-ended Questions","year":2019,"lang":"en","type":"dissertation","venue":"uO Research (University of Ottawa)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Mathematics education; Computer science; Closed-ended question; Information retrieval; Natural language processing; Psychology; Data science; Mathematics; Statistics","score_opus":0.1268726717033858,"score_gpt":0.3884350634991476,"score_spread":0.2615623917957618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3192525318","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32703322,0.0014294222,0.64254206,0.0010136769,0.00046035307,0.00061220076,0.003893735,0.017947437,0.0050678765],"genre_scores_gemma":[0.60997844,0.00042574032,0.36517608,0.00027467657,0.00031962956,0.00052366505,0.014772638,0.00035765555,0.008171493],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99438274,0.0020830594,0.0003484576,0.001677843,0.0011474874,0.00036035915],"domain_scores_gemma":[0.9874473,0.004617663,0.0012204868,0.0010481675,0.0050626723,0.0006037217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044223946,0.0014708408,0.0012031576,0.004695162,0.00048477118,0.002356216,0.0015620088,0.0026946687,0.0021607124],"category_scores_gemma":[0.014300117,0.00030111603,0.0012598411,0.0014646238,0.00044570593,0.0025617515,0.0019534095,0.0021071055,0.0058525046],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047426496,0.00084268063,0.030949034,0.0003784637,0.00012837091,0.00014450552,0.0005695743,0.007241339,0.040645633,0.0016087361,0.0130722355,0.90394515],"study_design_scores_gemma":[0.0000897429,0.0006984823,0.04126288,0.00015661697,0.0001512246,0.0003375711,0.0011564788,0.8660925,0.067508236,0.008063893,0.0143828755,0.00009951506],"about_ca_topic_score_codex":0.0014246785,"about_ca_topic_score_gemma":0.0014312875,"teacher_disagreement_score":0.004695162,"about_ca_system_score_codex":0.000525654,"about_ca_system_score_gemma":0.0010262781,"threshold_uncertainty_score":0.023388147},"labels":[],"label_agreement":null},{"id":"W3193068792","doi":"10.1162/coli_a_00422","title":"Probing Classifiers: Promises, Shortcomings, and Advances","year":2021,"lang":"en","type":"preprint","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation; Israel Science Foundation","keywords":"Computer science; Artificial intelligence; Variety (cybernetics); Classifier (UML); Machine learning; Property (philosophy); Artificial neural network; Deep neural networks; Natural language processing; Epistemology","score_opus":0.03569677429684543,"score_gpt":0.2887908698049079,"score_spread":0.2530940955080625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3193068792","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017404899,0.098570935,0.7594551,0.09758005,0.0027338502,0.00015598362,0.0005860307,0.0025234597,0.020989707],"genre_scores_gemma":[0.56065357,0.061848525,0.3392277,0.011915313,0.014773954,0.00057183433,0.0010943984,0.00100626,0.008908446],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98483,0.0076721157,0.00046026593,0.002329351,0.004201488,0.00050674705],"domain_scores_gemma":[0.90112007,0.0717757,0.0019394085,0.009236936,0.014194125,0.0017338203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03587117,0.0016751926,0.002665594,0.0032462147,0.0018990176,0.007862362,0.004325789,0.005503535,0.0037924557],"category_scores_gemma":[0.09453927,0.0010112543,0.0011449875,0.003755189,0.0048372424,0.024495792,0.0038425038,0.011816834,0.0032707432],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002787897,0.00013459774,0.0036708142,0.0005024605,0.00013540474,0.00008146407,0.00062821346,0.016644903,0.0009037099,0.46999562,0.045917664,0.4611063],"study_design_scores_gemma":[0.000024604618,0.00007951927,0.0006542942,0.00032224122,0.00005599769,0.0001294786,0.00023502858,0.17995542,0.0012066812,0.7693937,0.047880568,0.000062499734],"about_ca_topic_score_codex":0.0042032916,"about_ca_topic_score_gemma":0.001786568,"teacher_disagreement_score":0.03587117,"about_ca_system_score_codex":0.0033029024,"about_ca_system_score_gemma":0.00318854,"threshold_uncertainty_score":0.1897071},"labels":[],"label_agreement":null},{"id":"W3193215773","doi":"","title":"Detecting Interrogative Utterances with Recurrent Neural Networks","year":2015,"lang":"en","type":"article","venue":"Neural Information Processing Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Utterance; Computer science; Artificial neural network; Regularization (linguistics); Artificial intelligence; Recurrent neural network; Interrogative; Deep neural networks; Context (archaeology); Statement (logic); Speech recognition; Natural language processing; Machine learning; Linguistics","score_opus":0.04103528702438647,"score_gpt":0.2667657193270378,"score_spread":0.22573043230265133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3193215773","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.65195453,0.0025583205,0.3265243,0.0013735833,0.0003012958,0.0001858264,0.0024103564,0.005811172,0.00888053],"genre_scores_gemma":[0.9543119,0.0003030151,0.039839663,0.0001637343,0.0001309667,0.00007869757,0.0018818036,0.00013025274,0.0031598336],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986451,0.0005606495,0.00006126243,0.0003178903,0.00025682704,0.00015821734],"domain_scores_gemma":[0.9965693,0.0020607263,0.00040268968,0.0002753629,0.00059444783,0.00009746744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016198417,0.0011221734,0.00062128564,0.0012028185,0.00033361002,0.00097229285,0.0010544865,0.0010345401,0.0016691405],"category_scores_gemma":[0.0073263594,0.0003397318,0.0005675488,0.00067362504,0.0003771204,0.0016352572,0.0012047654,0.0012821546,0.0012045508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002628673,0.00070808467,0.051186103,0.0005849361,0.00059283385,0.0014971095,0.0032263047,0.07189059,0.11635477,0.0070029926,0.019738961,0.7245887],"study_design_scores_gemma":[0.000031699175,0.00018548733,0.017144775,0.0000598505,0.00013088294,0.00031902874,0.00081210973,0.941719,0.028674694,0.0059398366,0.0049146865,0.00006791218],"about_ca_topic_score_codex":0.0041628233,"about_ca_topic_score_gemma":0.0058526206,"teacher_disagreement_score":0.0041628233,"about_ca_system_score_codex":0.00064777356,"about_ca_system_score_gemma":0.00039952376,"threshold_uncertainty_score":0.008566678},"labels":[],"label_agreement":null},{"id":"W3194334866","doi":"10.1145/3469096.3474926","title":"MTLV","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Artificial intelligence; Multi-task learning; Task (project management); Machine learning; Generalization; Task analysis; Natural language processing","score_opus":0.027232784057806287,"score_gpt":0.24106642442919532,"score_spread":0.21383364037138902,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194334866","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00923731,0.0018407981,0.6305141,0.00063908467,0.00046466032,0.00039467582,0.032537185,0.30905104,0.015321128],"genre_scores_gemma":[0.14546733,0.0017462507,0.6530377,0.0012933408,0.00024451478,0.0022566032,0.14179157,0.018531954,0.035630733],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991479,0.00018594808,0.000073431096,0.00030469682,0.00018837482,0.00009948017],"domain_scores_gemma":[0.9988412,0.00042043172,0.00006743939,0.00038988533,0.00021607397,0.000064998174],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0013206374,0.0018860175,0.0007277615,0.0017236547,0.0006231629,0.0018747644,0.003120093,0.0018093133,0.04268675],"category_scores_gemma":[0.005772149,0.0008133445,0.0021880234,0.001671631,0.00032377956,0.002661686,0.0020079212,0.0022095807,0.028535347],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004990819,0.00029575636,0.0019288054,0.0010229722,0.00029293064,0.00023464314,0.00021500609,0.045597836,0.0056839557,0.014055855,0.30149215,0.62868106],"study_design_scores_gemma":[0.0002325501,0.00039152792,0.0015353342,0.00023004739,0.0001120567,0.00041365938,0.00010705237,0.73712784,0.021922922,0.04977577,0.1880493,0.00010196264],"about_ca_topic_score_codex":0.0071476605,"about_ca_topic_score_gemma":0.011668146,"teacher_disagreement_score":0.95731324,"about_ca_system_score_codex":0.0013514078,"about_ca_system_score_gemma":0.0015266765,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3194884147","doi":"10.1145/3447548.3470791","title":"Language Scaling","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Scaling; Taxonomy (biology); Pragmatics; Natural language processing; Natural language; Artificial intelligence; Data science; Linguistics","score_opus":0.01814573464516618,"score_gpt":0.2556594184063437,"score_spread":0.23751368376117754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194884147","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014070626,0.007884021,0.838007,0.013646345,0.002585558,0.0008336686,0.0044227294,0.013119575,0.10543054],"genre_scores_gemma":[0.27949676,0.015455553,0.5888163,0.011897637,0.004239729,0.0029089241,0.019612573,0.010515025,0.06705756],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99134606,0.0028143178,0.0006358039,0.0024612765,0.0021930726,0.0005494427],"domain_scores_gemma":[0.97662467,0.008347658,0.0008893654,0.007893822,0.005420049,0.00082452025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075373864,0.0016574118,0.0017203981,0.0031868885,0.0021112342,0.008016107,0.0044660578,0.0021245426,0.034079637],"category_scores_gemma":[0.040726755,0.000917581,0.0026340794,0.004136582,0.0023706474,0.019908724,0.009115966,0.004579656,0.022907557],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015604182,0.00021870794,0.002597703,0.0012702364,0.00016796464,0.00036812777,0.0015414574,0.0133709265,0.007016802,0.37656608,0.13472658,0.46199936],"study_design_scores_gemma":[0.000048372323,0.000092392555,0.0011458348,0.00044101317,0.00010128349,0.0007130177,0.0012751772,0.07820745,0.0057977247,0.4408008,0.47127625,0.00010066171],"about_ca_topic_score_codex":0.0030350839,"about_ca_topic_score_gemma":0.0020325128,"teacher_disagreement_score":0.034079637,"about_ca_system_score_codex":0.0026014266,"about_ca_system_score_gemma":0.003267485,"threshold_uncertainty_score":0.11400777},"labels":[],"label_agreement":null},{"id":"W3195010973","doi":"10.1145/3459637.3482011","title":"MS MARCO Chameleons: Challenging the MS MARCO Leaderboard with Extremely Obstinate Queries","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Microsoft (Canada); University of Waterloo","funders":"","keywords":"Computer science; Set (abstract data type); Task (project management); Information retrieval; Perspective (graphical); Artificial intelligence","score_opus":0.03365566842748198,"score_gpt":0.22261926057938417,"score_spread":0.18896359215190217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3195010973","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77316475,0.016509483,0.022320803,0.009595353,0.0028338043,0.0015781667,0.0985291,0.018912427,0.056556024],"genre_scores_gemma":[0.643014,0.002283831,0.08234582,0.0027614625,0.0012839071,0.00058484357,0.23507935,0.0013568492,0.031289857],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99695206,0.0007720973,0.00022830795,0.00068991503,0.0010402674,0.00031729206],"domain_scores_gemma":[0.99405944,0.0027908024,0.0004119436,0.0011742783,0.00086100434,0.00070254295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025327478,0.0018671536,0.0015887212,0.002457549,0.0017362119,0.0023570072,0.0019152034,0.0025985036,0.007369923],"category_scores_gemma":[0.013105017,0.00029989012,0.00088389445,0.0024211924,0.001169484,0.003568735,0.00212921,0.0022766232,0.0041803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004493756,0.0015978775,0.019718776,0.003035524,0.00038295955,0.0025771777,0.0013274198,0.027663015,0.016433375,0.0075916946,0.6358856,0.2792929],"study_design_scores_gemma":[0.0018008624,0.0035182892,0.0408888,0.0005435601,0.00029344828,0.005225515,0.007931583,0.34790128,0.05083969,0.015354058,0.52514887,0.0005540647],"about_ca_topic_score_codex":0.028154174,"about_ca_topic_score_gemma":0.061142143,"teacher_disagreement_score":0.028154174,"about_ca_system_score_codex":0.001700138,"about_ca_system_score_gemma":0.0019682755,"threshold_uncertainty_score":0.055980563},"labels":[],"label_agreement":null},{"id":"W3196070880","doi":"10.48550/arxiv.2111.05196","title":"NATURE: Natural Auxiliary Text Utterances for Realistic Spoken Language Evaluation","year":2021,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Utterance; Set (abstract data type); Spoken language; Simple (philosophy); Natural language; Natural language processing; Semantics (computer science); Natural language understanding; Artificial intelligence; Speech recognition; Programming language","score_opus":0.05376264222521476,"score_gpt":0.22166264853563356,"score_spread":0.1679000063104188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196070880","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41246554,0.0032221868,0.39793366,0.0017185374,0.0022820407,0.0020954146,0.06788743,0.0947464,0.017648818],"genre_scores_gemma":[0.6114268,0.00042167725,0.25052142,0.0005518675,0.00019501829,0.00205336,0.12629384,0.0028995287,0.0056365174],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943376,0.002919307,0.0004185191,0.0010306243,0.001053221,0.00024068073],"domain_scores_gemma":[0.99289495,0.0036415535,0.00032903417,0.0014994172,0.0013329939,0.00030201877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054666065,0.0025771912,0.0010670369,0.0012826346,0.00076689175,0.002104872,0.0027032194,0.0016758003,0.004769307],"category_scores_gemma":[0.019672818,0.00045643965,0.0008839774,0.0009493984,0.00072725466,0.0025895047,0.0022443058,0.002096136,0.004073781],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003730747,0.0022952715,0.015834125,0.00307655,0.000658106,0.000907794,0.0021392414,0.19133015,0.06336689,0.008461026,0.21153148,0.49666864],"study_design_scores_gemma":[0.0003335657,0.001109037,0.0074239364,0.00014460382,0.00011060124,0.0004392762,0.0010557661,0.8843183,0.053015906,0.010013186,0.041860756,0.00017497697],"about_ca_topic_score_codex":0.007594735,"about_ca_topic_score_gemma":0.01172482,"teacher_disagreement_score":0.007594735,"about_ca_system_score_codex":0.0012233369,"about_ca_system_score_gemma":0.0013132835,"threshold_uncertainty_score":0.028910518},"labels":[],"label_agreement":null},{"id":"W3197016261","doi":"10.14778/3476311.3476326","title":"CBench","year":2021,"lang":"en","type":"article","venue":"Proceedings of the VLDB Endowment","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Benchmarking; Suite; Benchmark (surveying); Question answering; Set (abstract data type); Task (project management); Quality (philosophy); Artificial intelligence; Information retrieval; Natural language processing; Programming language; Systems engineering","score_opus":0.01596362296791629,"score_gpt":0.2186539506523365,"score_spread":0.20269032768442022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197016261","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0424864,0.005157063,0.18930519,0.0015975416,0.0016657104,0.0021853617,0.08803118,0.54420537,0.12536617],"genre_scores_gemma":[0.2062019,0.0026297627,0.288634,0.0021423271,0.00032669737,0.00365202,0.37554494,0.057456892,0.06341153],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99283785,0.0015303052,0.0008526623,0.0011998389,0.0029685453,0.000610826],"domain_scores_gemma":[0.9865362,0.003396722,0.00050821406,0.0039796494,0.0050370153,0.0005423375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004941433,0.0022834754,0.0012854395,0.0044938484,0.0013302249,0.003980914,0.0041955146,0.0019064222,0.033811454],"category_scores_gemma":[0.020909846,0.0008568113,0.001444134,0.0044206884,0.00084607943,0.0048653595,0.0042488193,0.0025178492,0.031983517],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016439591,0.00090055936,0.0055253743,0.002940087,0.00022593599,0.0003418904,0.00091869925,0.010772586,0.012347705,0.027381452,0.64675266,0.29024905],"study_design_scores_gemma":[0.00042149136,0.00071544683,0.006729153,0.00050473673,0.00012825176,0.00051454996,0.0007034084,0.09204452,0.02749698,0.034175634,0.8363254,0.00024050692],"about_ca_topic_score_codex":0.009404302,"about_ca_topic_score_gemma":0.008403559,"teacher_disagreement_score":0.033811454,"about_ca_system_score_codex":0.001617888,"about_ca_system_score_gemma":0.0029009045,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W3197275000","doi":"10.18653/v1/2021.emnlp-main.181","title":"Unsupervised Conversation Disentanglement through Co-Training","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Conversation; Computer science; Classifier (UML); Session (web analytics); Artificial intelligence; Machine learning; Reinforcement learning; Training set; Speech recognition; World Wide Web","score_opus":0.11382687623688474,"score_gpt":0.4306470469350119,"score_spread":0.3168201706981271,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197275000","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.066704236,0.0010808976,0.92131114,0.00038291048,0.00015653625,0.00024587044,0.0004833981,0.0065861526,0.003048972],"genre_scores_gemma":[0.72213465,0.0003271153,0.26307282,0.0004192645,0.0002217147,0.0005777828,0.0028105723,0.0006790292,0.009757007],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972361,0.0009276456,0.00009624664,0.0012210725,0.00028217948,0.00023667744],"domain_scores_gemma":[0.9951775,0.002820145,0.00036513578,0.0006993425,0.0005706384,0.00036717468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023654432,0.0023565218,0.0014271357,0.0016353289,0.0008437277,0.0010245763,0.0025154206,0.0015668441,0.003083415],"category_scores_gemma":[0.00864773,0.00075415656,0.0012852412,0.0010280003,0.0008892327,0.0036741951,0.0031955193,0.0040841578,0.002427146],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013524259,0.0008854615,0.009288321,0.00040053,0.00027183103,0.00029788818,0.0014338327,0.08755701,0.03732668,0.0038808489,0.010366512,0.84693867],"study_design_scores_gemma":[0.000027792967,0.00011357133,0.0013784858,0.000022860042,0.000039083436,0.00009037356,0.00015422136,0.9798365,0.009430743,0.0056388215,0.0032372598,0.000030272795],"about_ca_topic_score_codex":0.00438437,"about_ca_topic_score_gemma":0.00933098,"teacher_disagreement_score":0.00438437,"about_ca_system_score_codex":0.000758416,"about_ca_system_score_gemma":0.0015311359,"threshold_uncertainty_score":0.012509763},"labels":[],"label_agreement":null},{"id":"W3197317199","doi":"10.18653/v1/2022.starsem-1.16","title":"A Generative Approach for Mitigating Structural Biases in Natural Language Inference","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation","keywords":"Discriminative model; Generative grammar; Computer science; Generative model; Inference; Artificial intelligence; Task (project management); Machine learning; Context (archaeology); Pattern recognition (psychology)","score_opus":0.06508462960440754,"score_gpt":0.32992578380389664,"score_spread":0.2648411541994891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197317199","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010969961,0.0002762296,0.9859779,0.0004560395,0.00002956396,0.000061962724,0.0001030495,0.0010591847,0.0010661224],"genre_scores_gemma":[0.6067366,0.00057737535,0.38467252,0.0016402983,0.00032806152,0.0004996224,0.0010902053,0.0008494901,0.0036057632],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994582,0.0030182349,0.0002135483,0.001018108,0.00089038647,0.00027775185],"domain_scores_gemma":[0.9788406,0.014692819,0.0008859103,0.00431903,0.0009192094,0.0003424608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008525584,0.0015600541,0.0016161276,0.0015054037,0.0010461456,0.0017376933,0.0031964125,0.0022341912,0.0027209865],"category_scores_gemma":[0.033124946,0.0014796646,0.0017987936,0.0014010833,0.0027801504,0.003930483,0.0046325778,0.0047674514,0.0011350076],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042420343,0.0003194586,0.009364375,0.00035033937,0.0004085423,0.0003381042,0.0010644903,0.59065,0.014665134,0.18688175,0.007067489,0.18846603],"study_design_scores_gemma":[0.000032086093,0.00004705707,0.00041021346,0.000027632363,0.000043080847,0.000107544925,0.00002054599,0.9068085,0.0029079074,0.08799616,0.001577331,0.00002189693],"about_ca_topic_score_codex":0.0031107217,"about_ca_topic_score_gemma":0.007249512,"teacher_disagreement_score":0.008525584,"about_ca_system_score_codex":0.0017320977,"about_ca_system_score_gemma":0.002179517,"threshold_uncertainty_score":0.045088112},"labels":[],"label_agreement":null},{"id":"W3197492799","doi":"10.1007/978-3-030-86331-9_23","title":"Sparse Document Analysis Using Beta-Liouville Naive Bayes with Vocabulary Knowledge","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; BETA (programming language); Bayes' theorem; Artificial intelligence; Naive Bayes classifier; Vocabulary; Natural language processing; Machine learning; Bayesian probability; Programming language; Philosophy; Linguistics; Support vector machine","score_opus":0.024601828288167968,"score_gpt":0.260545632630907,"score_spread":0.23594380434273904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197492799","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037090061,0.00094353297,0.99325824,0.00014859266,0.0000891207,0.00006927039,0.0002000395,0.00094092934,0.0006414213],"genre_scores_gemma":[0.13780129,0.0015003536,0.8490766,0.00034602813,0.00071907614,0.00036629837,0.003163549,0.00036910173,0.006657627],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957832,0.0016085982,0.00043760042,0.00075788744,0.0011630093,0.00024981727],"domain_scores_gemma":[0.99134314,0.0060812193,0.00026014124,0.00078516686,0.0013823701,0.00014804023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004192374,0.0011872443,0.002759842,0.004696221,0.0015837958,0.0035462405,0.002511479,0.0021581212,0.0045507643],"category_scores_gemma":[0.01470317,0.0010593559,0.0024085015,0.004210253,0.001022413,0.0046637277,0.0018697467,0.0029318612,0.0034778095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037782817,0.00029941447,0.0020025421,0.00041512653,0.00031887917,0.0002030764,0.00021621556,0.07040436,0.006690847,0.020693736,0.012470875,0.88590705],"study_design_scores_gemma":[0.00003442281,0.000050761326,0.000338792,0.000044496835,0.00009746849,0.00014660499,0.000046232293,0.95328397,0.001992342,0.04112548,0.002809176,0.00003031719],"about_ca_topic_score_codex":0.0076657548,"about_ca_topic_score_gemma":0.010204676,"teacher_disagreement_score":0.0076657548,"about_ca_system_score_codex":0.0011113113,"about_ca_system_score_gemma":0.002062989,"threshold_uncertainty_score":0.022171617},"labels":[],"label_agreement":null},{"id":"W3198035914","doi":"10.1007/978-3-030-86337-1_16","title":"A More Effective Sentence-Wise Text Segmentation Approach Using BERT","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Segmentation; Encoder; Artificial intelligence; Sentence; Transformer; Natural language processing; Text segmentation; Feature (linguistics); Heuristic; Feature engineering; Language model; Deep learning; Natural language","score_opus":0.025975732776636576,"score_gpt":0.2658582777646642,"score_spread":0.23988254498802764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3198035914","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013991864,0.0012211634,0.94198275,0.0006212687,0.00088178023,0.00042919288,0.0026484118,0.03243456,0.005788951],"genre_scores_gemma":[0.086761564,0.00059058965,0.8768255,0.000481641,0.0006446001,0.00032871452,0.010658698,0.0029737167,0.020734975],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99857914,0.00023637299,0.00012545595,0.0005692358,0.00034622557,0.00014367446],"domain_scores_gemma":[0.99833906,0.0005368031,0.00006616305,0.00026332794,0.0006698029,0.00012477879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009910172,0.0025640042,0.0022027432,0.0038612173,0.0015312661,0.0028690938,0.0019849562,0.002099946,0.029174848],"category_scores_gemma":[0.002229796,0.00084404665,0.0019072864,0.003094906,0.000429223,0.0036223873,0.0021190916,0.0022125063,0.022647647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006898867,0.00023006542,0.00047798952,0.00047201588,0.00013535855,0.0002605633,0.00022731548,0.0037324065,0.1719147,0.0038991405,0.03439333,0.78356725],"study_design_scores_gemma":[0.0002106429,0.0005844026,0.0036582125,0.000110094246,0.00046276362,0.0011267039,0.0007858861,0.70961165,0.17251363,0.014352307,0.0963772,0.0002065429],"about_ca_topic_score_codex":0.005765282,"about_ca_topic_score_gemma":0.0091684265,"teacher_disagreement_score":0.029174848,"about_ca_system_score_codex":0.00058919337,"about_ca_system_score_gemma":0.0016207905,"threshold_uncertainty_score":0.097599566},"labels":[],"label_agreement":null},{"id":"W3198536471","doi":"10.1145/3446426","title":"Multi-Stage Conversational Passage Retrieval: An Approach to Fusing Term Importance Estimation and Neural Query Rewriting","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Canada First Research Excellence Fund; Natural Sciences and Engineering Research Council of Canada; Ministry of Science and Technology, Taiwan","keywords":"Computer science; Information retrieval; Leverage (statistics); Query expansion; Term (time); Artificial intelligence; Query language; Rewriting; Context (archaeology); Natural language processing; Programming language","score_opus":0.059295514660712485,"score_gpt":0.2863921894774662,"score_spread":0.22709667481675372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3198536471","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009907813,0.0011075189,0.986388,0.00021220038,0.00008140452,0.0001555261,0.000089408095,0.0011877074,0.0008704045],"genre_scores_gemma":[0.31503895,0.0009658606,0.6757134,0.00036382276,0.0005882493,0.00035267955,0.0007313958,0.00035485905,0.005890726],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973309,0.0010319696,0.00020964534,0.0007158914,0.00054034137,0.00017125657],"domain_scores_gemma":[0.99725467,0.0013373437,0.00022849032,0.00043220178,0.00061931304,0.00012805854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031429315,0.0013730623,0.0016636582,0.002006054,0.0006960273,0.0013361458,0.0028554206,0.00166488,0.0021140913],"category_scores_gemma":[0.007915975,0.00065158855,0.0016520473,0.0016284225,0.00085809367,0.002709312,0.0017679031,0.0016984459,0.0011735277],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054965325,0.00048761757,0.0014806816,0.000550738,0.00031636452,0.0005804274,0.0012202966,0.10987189,0.08651904,0.018181084,0.0069357227,0.77330655],"study_design_scores_gemma":[0.000020516462,0.00020027148,0.0006005541,0.000015089794,0.00011307302,0.00018455891,0.00008721006,0.97237,0.013845848,0.008039986,0.0044635623,0.000059237016],"about_ca_topic_score_codex":0.0076600434,"about_ca_topic_score_gemma":0.0075337593,"teacher_disagreement_score":0.0076600434,"about_ca_system_score_codex":0.0009827052,"about_ca_system_score_gemma":0.0014371332,"threshold_uncertainty_score":0.01662165},"labels":[],"label_agreement":null},{"id":"W3198561049","doi":"10.18653/v1/2021.findings-emnlp.84","title":"When Retriever-Reader Meets Scenario-Based Multiple-Choice Questions","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Variety (cybernetics); Relevance (law); Artificial intelligence; Reading (process); Information retrieval; Natural language processing; Word (group theory); Weighting; Question answering; Noise (video); Labrador Retriever; Machine learning; Linguistics","score_opus":0.051033749843265454,"score_gpt":0.280786606815107,"score_spread":0.22975285697184153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3198561049","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38134515,0.00328234,0.5696626,0.006812482,0.00042902172,0.0010814088,0.0068234345,0.015623313,0.014940292],"genre_scores_gemma":[0.85776967,0.00051511126,0.1133314,0.001704688,0.00040995446,0.00037677126,0.0155276675,0.0007330146,0.009631676],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9858284,0.006973789,0.0009030668,0.0044134865,0.0012195169,0.0006618213],"domain_scores_gemma":[0.96368176,0.024346361,0.0022475598,0.0060151685,0.0025707355,0.0011384567],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014135405,0.0016977271,0.002010082,0.0017717741,0.0009842583,0.003812331,0.0031710847,0.007878104,0.005420117],"category_scores_gemma":[0.05872941,0.0010240712,0.0017248426,0.0014385025,0.001328808,0.016424587,0.0034348625,0.0033913995,0.00709948],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005610919,0.002897007,0.13538851,0.0038819634,0.0012181277,0.0054258252,0.013155322,0.13650353,0.08929797,0.055215437,0.119622305,0.4317831],"study_design_scores_gemma":[0.00020040789,0.00071325246,0.013196185,0.00009270908,0.00018758654,0.002166213,0.002020771,0.87037253,0.022316867,0.056554876,0.032016937,0.00016165605],"about_ca_topic_score_codex":0.004046728,"about_ca_topic_score_gemma":0.004288603,"teacher_disagreement_score":0.014135405,"about_ca_system_score_codex":0.0013855472,"about_ca_system_score_gemma":0.0010766914,"threshold_uncertainty_score":0.074756086},"labels":[],"label_agreement":null},{"id":"W3198948523","doi":"10.18653/v1/2021.findings-emnlp.90","title":"Exploring Decomposition for Table-based Fact Verification","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Benchmark (surveying); Decomposition; Table (database); Parsing; Simple (philosophy); Artificial intelligence; Programming language; Natural language; Machine learning; Formal verification; Natural language processing; Theoretical computer science; Data mining","score_opus":0.21935186610806487,"score_gpt":0.3194389970665212,"score_spread":0.10008713095845631,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3198948523","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05569416,0.0004984189,0.92419386,0.000919083,0.000070595706,0.00024108584,0.003766848,0.013164743,0.0014511985],"genre_scores_gemma":[0.4072933,0.00036248344,0.56961006,0.00032327202,0.00008649706,0.00028743898,0.019468954,0.0009210691,0.0016469602],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979234,0.00061053707,0.00017359448,0.0006303945,0.00048611127,0.00017597037],"domain_scores_gemma":[0.99204165,0.0050735082,0.00056374434,0.001404266,0.0007231838,0.00019363951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002351296,0.0011494884,0.00084108487,0.0019862084,0.0006811179,0.0020398458,0.0015499457,0.001115221,0.004651885],"category_scores_gemma":[0.013801905,0.0006216084,0.0023577653,0.0015707308,0.0010213587,0.006097562,0.0025570786,0.0022882104,0.0018602305],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009788891,0.0005321973,0.021422166,0.0013304335,0.00028131763,0.0011252586,0.0013979739,0.21307145,0.03574247,0.06977389,0.029068647,0.6252753],"study_design_scores_gemma":[0.0000370401,0.00007719776,0.0009898242,0.000066369445,0.000048847076,0.00018427502,0.00017033295,0.91184247,0.012079193,0.06675914,0.0077234465,0.000021937038],"about_ca_topic_score_codex":0.00493911,"about_ca_topic_score_gemma":0.008531386,"teacher_disagreement_score":0.00493911,"about_ca_system_score_codex":0.0012427936,"about_ca_system_score_gemma":0.0027408116,"threshold_uncertainty_score":0.0155620575},"labels":[],"label_agreement":null},{"id":"W3199203364","doi":"10.18653/v1/2021.findings-emnlp.326","title":"Uncovering Implicit Gender Bias in Narratives through Commonsense Inference","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Narrative; Focus (optics); Inference; Commonsense reasoning; Commonsense knowledge; Implicit bias; Intellect; Cognitive psychology; Computer science; Space (punctuation); Psychology; Artificial intelligence; Natural language processing; Epistemology; Social psychology; Linguistics","score_opus":0.16759340898275887,"score_gpt":0.3425471001759142,"score_spread":0.17495369119315535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3199203364","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38767368,0.00047667473,0.5968089,0.0016073833,0.000109116874,0.00023266724,0.0028646397,0.002222979,0.008003954],"genre_scores_gemma":[0.86949044,0.000116915355,0.12647708,0.00014542058,0.000034205626,0.000117457486,0.0025049273,0.00020322262,0.00091036584],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99541557,0.002711364,0.0002012187,0.0008556645,0.00063381885,0.00018230522],"domain_scores_gemma":[0.9575174,0.036536727,0.001609583,0.0030192998,0.0010425334,0.00027451295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057518664,0.0008378868,0.00040883254,0.001269636,0.00058000936,0.0022662263,0.00132035,0.00097612705,0.00268595],"category_scores_gemma":[0.037315134,0.00054808805,0.00077572465,0.00070097967,0.0013241374,0.004045404,0.0024103904,0.0017517254,0.0008455568],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011430187,0.00055293995,0.09891938,0.0018913224,0.00048506868,0.0025797328,0.037142865,0.12682451,0.047649752,0.19318254,0.015127505,0.4745014],"study_design_scores_gemma":[0.00007581041,0.00009269699,0.007496783,0.00021057042,0.00009594664,0.0006131214,0.0026013572,0.7937011,0.023758547,0.15573458,0.015558217,0.00006124128],"about_ca_topic_score_codex":0.0013737326,"about_ca_topic_score_gemma":0.002751682,"teacher_disagreement_score":0.0057518664,"about_ca_system_score_codex":0.001074273,"about_ca_system_score_gemma":0.0008336083,"threshold_uncertainty_score":0.030419111},"labels":[],"label_agreement":null},{"id":"W3199446243","doi":"10.1109/tnnls.2021.3112045","title":"Combining Knowledge Graph and Word Embeddings for Spherical Topic Modeling","year":2021,"lang":"en","type":"article","venue":"IEEE Transactions on Neural Networks and Learning Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Interpretability; Hypersphere; Artificial intelligence; Natural language processing; Topic model; Leverage (statistics); Discriminative model; Word (group theory); Probabilistic logic; Graph; Set (abstract data type); Machine learning; Theoretical computer science; Mathematics","score_opus":0.02569691609597778,"score_gpt":0.25348402159658917,"score_spread":0.22778710550061138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3199446243","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012657456,0.00065416907,0.9847749,0.00023323992,0.000031895546,0.000043725457,0.00037060367,0.00050304737,0.0007309474],"genre_scores_gemma":[0.60795474,0.0029068762,0.3768208,0.0004957289,0.00030437345,0.0004696982,0.005793536,0.0004767312,0.004777528],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99891245,0.00038705594,0.00007299087,0.00037251928,0.00018249956,0.00007237936],"domain_scores_gemma":[0.9979278,0.0011778601,0.00024557128,0.00030467904,0.00026695203,0.000077202756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012544839,0.0012306707,0.000950423,0.0028773532,0.00040623025,0.0012519492,0.0014326703,0.001325182,0.0014219654],"category_scores_gemma":[0.006194836,0.0005116735,0.0015922346,0.0039434754,0.0008794696,0.0041914927,0.0015455253,0.0018357699,0.0012286631],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002952173,0.00020518171,0.004849191,0.0004255845,0.00031624959,0.00032298753,0.0007883517,0.46747115,0.00871652,0.10438528,0.0099866185,0.40223762],"study_design_scores_gemma":[0.000010461616,0.000028407234,0.0004768977,0.000015564572,0.000028868784,0.00006973316,0.000036612328,0.94948155,0.00059794326,0.04722265,0.0020134535,0.000017804736],"about_ca_topic_score_codex":0.005893331,"about_ca_topic_score_gemma":0.009244579,"teacher_disagreement_score":0.005893331,"about_ca_system_score_codex":0.0008959805,"about_ca_system_score_gemma":0.00091897434,"threshold_uncertainty_score":0.011718035},"labels":[],"label_agreement":null},{"id":"W3200130628","doi":"10.18653/v1/2021.emnlp-main.122","title":"Conditional probing: measuring usable information beyond a baseline","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Baseline (sea); USable; Representation (politics); Computer science; Word (group theory); Property (philosophy); Identity (music); Natural language processing; Artificial intelligence; Speech recognition; Mathematics","score_opus":0.05992423051901477,"score_gpt":0.36906375505046807,"score_spread":0.3091395245314533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3200130628","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34397712,0.0006867602,0.64853764,0.00076061365,0.00005536628,0.0001451458,0.0010068103,0.0014616976,0.0033688403],"genre_scores_gemma":[0.96680796,0.00013689509,0.031435136,0.00019942274,0.000035210716,0.00012453234,0.00075925933,0.00016965819,0.000331786],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949567,0.0022296172,0.000307234,0.0013680046,0.00082643016,0.00031208413],"domain_scores_gemma":[0.92483455,0.057534855,0.004006672,0.010535407,0.0019197218,0.0011688009],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008933778,0.0013249208,0.0012642334,0.0013374424,0.00064822217,0.0023125778,0.0016508137,0.002263468,0.0021946274],"category_scores_gemma":[0.06834812,0.0007824252,0.0009894497,0.0014457246,0.00306382,0.010359012,0.0037675377,0.004318867,0.00036004937],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0061236657,0.0013889676,0.10784878,0.0018218855,0.0016600897,0.00062722474,0.0028148182,0.30286637,0.10793937,0.1411985,0.0052264095,0.32048392],"study_design_scores_gemma":[0.00009316794,0.0013932255,0.02475498,0.00014073051,0.00038847493,0.00039138738,0.0003171033,0.6270015,0.036651038,0.306761,0.0019212407,0.00018602631],"about_ca_topic_score_codex":0.0011137223,"about_ca_topic_score_gemma":0.000939293,"teacher_disagreement_score":0.008933778,"about_ca_system_score_codex":0.001101396,"about_ca_system_score_gemma":0.0007890836,"threshold_uncertainty_score":0.047246933},"labels":[],"label_agreement":null},{"id":"W3200439183","doi":"10.1609/aaai.v36i10.21322","title":"BROS: A Pre-trained Language Model Focusing on Text and Layout for Better Key Information Extraction from Documents","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":120,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Key (lock); Natural language processing; Task (project management); Space (punctuation); Artificial intelligence; Masking (illustration); Information extraction; Language model; Semantics (computer science); Information retrieval","score_opus":0.048441363007867104,"score_gpt":0.30151575368776345,"score_spread":0.25307439067989634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3200439183","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043104798,0.0021887997,0.8723886,0.00068919855,0.00078602467,0.0005605002,0.005868124,0.06819866,0.006215296],"genre_scores_gemma":[0.22315226,0.0010481264,0.7164442,0.0010954649,0.00030570364,0.0008425877,0.031557124,0.0015298598,0.024024662],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99923694,0.00013983996,0.000054848515,0.0003145583,0.00015370181,0.00010027124],"domain_scores_gemma":[0.9984743,0.0005983987,0.0000996888,0.00029208214,0.00043663575,0.00009887557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013993804,0.0024418077,0.0012953564,0.0025173642,0.00048492948,0.0013910056,0.0030158896,0.0017846104,0.0070207887],"category_scores_gemma":[0.002824639,0.0007798419,0.0017080479,0.0013679581,0.00053632597,0.003953326,0.0016076664,0.002628723,0.007925421],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007048042,0.00047096546,0.0018206452,0.0007721414,0.00030838017,0.000304377,0.00018479241,0.05006388,0.051830184,0.0038378895,0.06878596,0.82091606],"study_design_scores_gemma":[0.00012734489,0.0002765126,0.0010898325,0.000057267785,0.000093224844,0.0002668659,0.00013018564,0.9455866,0.03200076,0.0027934627,0.017493302,0.0000846174],"about_ca_topic_score_codex":0.01229339,"about_ca_topic_score_gemma":0.022686183,"teacher_disagreement_score":0.01229339,"about_ca_system_score_codex":0.0008689368,"about_ca_system_score_gemma":0.0022086063,"threshold_uncertainty_score":0.024443686},"labels":[],"label_agreement":null},{"id":"W3200660033","doi":"10.1145/3459637.3482159","title":"Predicting Efficiency/Effectiveness Trade-offs for Dense vs. Sparse Retrieval Strategy Selection","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Leverage (statistics); Embedding; Classifier (UML); Artificial intelligence","score_opus":0.038431037457589544,"score_gpt":0.28088003111487775,"score_spread":0.2424489936572882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3200660033","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.746518,0.013069501,0.22478653,0.0016354077,0.00013693857,0.0006097508,0.0016669413,0.003862241,0.0077147093],"genre_scores_gemma":[0.9421868,0.0014848149,0.05200191,0.0002875548,0.00014366547,0.00017641654,0.0018064019,0.00018998889,0.0017224532],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976121,0.00056467275,0.00028779393,0.0005334548,0.0006834903,0.00031847288],"domain_scores_gemma":[0.98887247,0.00870396,0.0006141895,0.0007412201,0.00081520685,0.00025306898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034957395,0.0011001608,0.0016182985,0.0022612948,0.000375614,0.0015037655,0.0010193208,0.0014494136,0.001477198],"category_scores_gemma":[0.019999806,0.00034056228,0.0006276222,0.001708443,0.00070972485,0.0030211865,0.0005696217,0.00094720017,0.0014372119],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026762867,0.0014884546,0.04930082,0.0012983697,0.0004295679,0.0004961125,0.00023072893,0.23935504,0.034014042,0.004989871,0.020473048,0.64524764],"study_design_scores_gemma":[0.000120073804,0.0006407763,0.0047254604,0.000026609106,0.00012813705,0.00037905807,0.00010193379,0.97858137,0.009673897,0.0041336906,0.0014524547,0.00003647374],"about_ca_topic_score_codex":0.003997492,"about_ca_topic_score_gemma":0.004892233,"teacher_disagreement_score":0.003997492,"about_ca_system_score_codex":0.0009137,"about_ca_system_score_gemma":0.0010035723,"threshold_uncertainty_score":0.018487453},"labels":[],"label_agreement":null},{"id":"W3200891190","doi":"","title":"Inspecting the Factuality of Hallucinated Entities in Abstractive Summarization.","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Hallucinating; Automatic summarization; Computer science; Natural language processing; Artificial intelligence; Information retrieval","score_opus":0.09452848709654141,"score_gpt":0.1948844294690246,"score_spread":0.10035594237248319,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3200891190","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12108992,0.0031337554,0.86223096,0.00087106315,0.00026553692,0.0002510289,0.0023701128,0.0062318034,0.0035557763],"genre_scores_gemma":[0.71467996,0.0011744055,0.27267927,0.00018828965,0.00034199218,0.00013006756,0.0064391675,0.00038900904,0.0039778682],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990583,0.00032190533,0.00008754721,0.00023291529,0.00024234792,0.000057010562],"domain_scores_gemma":[0.99129325,0.0051985784,0.0010846481,0.00087156665,0.001356119,0.00019580171],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021036782,0.0010658554,0.00049066346,0.002334013,0.00045579637,0.001601113,0.0008410392,0.00088783336,0.0015325026],"category_scores_gemma":[0.014002127,0.00030408867,0.00051748526,0.0010587273,0.00042557777,0.0026049525,0.0011526599,0.0009963129,0.0011814445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008475952,0.00015245611,0.015123369,0.0014891246,0.00036339654,0.00086456357,0.002591572,0.033194534,0.11081934,0.007821875,0.020589557,0.8061426],"study_design_scores_gemma":[0.00008119185,0.0007137981,0.028251342,0.0002313788,0.00052713853,0.0013255306,0.0015805552,0.76934874,0.12870927,0.031555258,0.03751811,0.00015770119],"about_ca_topic_score_codex":0.0015183042,"about_ca_topic_score_gemma":0.0031080744,"teacher_disagreement_score":0.002334013,"about_ca_system_score_codex":0.00040909022,"about_ca_system_score_gemma":0.00042704056,"threshold_uncertainty_score":0.011125386},"labels":[],"label_agreement":null},{"id":"W3201341200","doi":"10.18653/v1/2021.emnlp-main.42","title":"Mitigating Language-Dependent Ethnic Bias in BERT","year":2021,"lang":"en","type":"preprint","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Ministry of Science and ICT, South Korea; National Research Foundation of Korea; National Research Foundation","keywords":"Categorical variable; Computer science; Ethnic group; German; Turkish; Natural language processing; Metric (unit); Language model; Linguistics; Arabic; Gender bias; Word (group theory); Artificial intelligence; Psychology; Machine learning; Sociology; Social psychology","score_opus":0.13897384373430055,"score_gpt":0.44294323317028,"score_spread":0.30396938943597945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3201341200","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7023168,0.0005701412,0.2899924,0.00050632353,0.000063073225,0.000077681565,0.00042014185,0.0015619858,0.004491512],"genre_scores_gemma":[0.9588619,0.00009810258,0.0386928,0.00008421125,0.00003480532,0.000037761907,0.0007932257,0.00020682752,0.0011903669],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99821115,0.0009891189,0.000097777294,0.00026438167,0.00028027841,0.00015735933],"domain_scores_gemma":[0.9901862,0.0063448614,0.0008828929,0.00113197,0.0012325338,0.00022150694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042722323,0.0008389969,0.0006941632,0.0010942676,0.0008362728,0.0011044266,0.00074633333,0.00065994536,0.001094152],"category_scores_gemma":[0.016233912,0.00031041572,0.0005035033,0.0011461006,0.000544026,0.0024548722,0.0017495163,0.0010994577,0.0005672492],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010598901,0.00031177868,0.18833457,0.00038015627,0.0004619137,0.0006083252,0.0036747633,0.33554143,0.019553743,0.024522182,0.007216659,0.41833454],"study_design_scores_gemma":[0.000028000533,0.000093483875,0.020144563,0.000033033586,0.000061063496,0.00024746856,0.0012205441,0.9478813,0.0074704774,0.018794637,0.0039735045,0.00005197468],"about_ca_topic_score_codex":0.007853147,"about_ca_topic_score_gemma":0.017671999,"teacher_disagreement_score":0.007853147,"about_ca_system_score_codex":0.000648275,"about_ca_system_score_gemma":0.00094556547,"threshold_uncertainty_score":0.022594035},"labels":[],"label_agreement":null},{"id":"W3201759467","doi":"","title":"TopiOCQA: Open-domain Conversational Question Answeringwith Topic Switching.","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Conversation; Computer science; Question answering; Domain (mathematical analysis); Natural language processing; Open domain; Interdependence; Artificial intelligence; Information retrieval; Code (set theory); Relevance (law); Limiting; Linguistics; Set (abstract data type); Programming language","score_opus":0.06853225001111356,"score_gpt":0.20295636979498613,"score_spread":0.1344241197838726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3201759467","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08935883,0.008804669,0.112678654,0.0029464916,0.00120249,0.004454778,0.65932614,0.095718935,0.025509039],"genre_scores_gemma":[0.104988724,0.000710167,0.11503844,0.0011697107,0.00022516002,0.002988358,0.7666423,0.0009872145,0.007249975],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9955948,0.001751579,0.0003475764,0.0012928718,0.00070128805,0.00031178733],"domain_scores_gemma":[0.99269694,0.003338097,0.00042606724,0.0017140185,0.0012240021,0.0006008712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030797322,0.0027417631,0.0012471825,0.0038045612,0.0018366614,0.002523021,0.0040443824,0.003319983,0.008638998],"category_scores_gemma":[0.019119166,0.00064711634,0.0016424255,0.0027248738,0.0008515583,0.0052985162,0.0052498668,0.003058212,0.009745427],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016233462,0.0010933584,0.008567476,0.0054815714,0.00044235744,0.0006368675,0.0023811324,0.010558476,0.013348737,0.006655249,0.79402745,0.15518394],"study_design_scores_gemma":[0.0010088471,0.00081041426,0.022397574,0.000711528,0.00033041288,0.0018011499,0.0029909634,0.28997353,0.02822966,0.026670078,0.62465936,0.00041647145],"about_ca_topic_score_codex":0.03139611,"about_ca_topic_score_gemma":0.048560787,"teacher_disagreement_score":0.03139611,"about_ca_system_score_codex":0.0022714736,"about_ca_system_score_gemma":0.0032434296,"threshold_uncertainty_score":0.062426746},"labels":[],"label_agreement":null},{"id":"W3201973725","doi":"10.18280/ria.350404","title":"Query-Based Retrieval Using Universal Sentence Encoder","year":2021,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Sentence; Computer science; Encoder; Natural language processing; Word (group theory); Artificial intelligence; Context (archaeology); Embedding; Word embedding; Question answering; Linguistics","score_opus":0.06598875846270742,"score_gpt":0.28072443677427783,"score_spread":0.21473567831157042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3201973725","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09674701,0.003010772,0.8522914,0.0007384405,0.00031769363,0.0005274367,0.006020427,0.033380512,0.0069662794],"genre_scores_gemma":[0.6196706,0.0010526851,0.34713858,0.00073244306,0.000211929,0.00049667124,0.018937899,0.0005708751,0.011188295],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994406,0.00015866716,0.00005139941,0.00018075104,0.000112154376,0.000056481793],"domain_scores_gemma":[0.99914587,0.00029760925,0.000050035498,0.00019183075,0.0002803772,0.00003426293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00093947025,0.0008161584,0.0007510504,0.0010893418,0.00027770616,0.0005894953,0.0008607383,0.0006202417,0.0043022935],"category_scores_gemma":[0.0027029123,0.00028981362,0.00067018415,0.00080509065,0.00033839807,0.0027043996,0.0010066796,0.000946281,0.002872681],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006393668,0.0004523761,0.0030986299,0.00065391127,0.00016567323,0.0002803004,0.00040618563,0.03861735,0.03992675,0.013228427,0.046904344,0.85562664],"study_design_scores_gemma":[0.0000983422,0.00036556335,0.0015892378,0.00003984096,0.0001238483,0.0003574433,0.00013865075,0.9331362,0.030390628,0.018613305,0.015100882,0.00004615866],"about_ca_topic_score_codex":0.006156527,"about_ca_topic_score_gemma":0.009571591,"teacher_disagreement_score":0.006156527,"about_ca_system_score_codex":0.00078068767,"about_ca_system_score_gemma":0.0009784148,"threshold_uncertainty_score":0.014392614},"labels":[],"label_agreement":null},{"id":"W3202137189","doi":"10.48550/arxiv.2109.13066","title":"Prefix-to-SQL: Text-to-SQL Generation from Incomplete User Questions","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; SQL; Prefix; Task (project management); Data definition language; Benchmark (surveying); Construct (python library); SQL injection; Stored procedure; Metric (unit); SQL/PSM; Programming language; Database; Query by Example; Information retrieval","score_opus":0.11416704476597267,"score_gpt":0.20239458114336342,"score_spread":0.08822753637739075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3202137189","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14376448,0.0015894485,0.5748853,0.0017859432,0.00068517553,0.0019270029,0.030843826,0.2334489,0.011069847],"genre_scores_gemma":[0.39624593,0.00044115508,0.49632606,0.0011453805,0.0002093455,0.0015224778,0.09047975,0.004656579,0.008973342],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99704033,0.0011713742,0.00028219484,0.0008054112,0.0005400176,0.00016075146],"domain_scores_gemma":[0.9886148,0.0071918066,0.00043012644,0.0021448408,0.0011742648,0.00044417242],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027662735,0.0024308502,0.0010062016,0.0016613418,0.0005066656,0.0017370931,0.0027443185,0.0023353659,0.014973057],"category_scores_gemma":[0.017931808,0.0005588988,0.0013516762,0.0010596123,0.0006883027,0.0042555244,0.0030918356,0.0019516764,0.007855507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022372806,0.001356586,0.017825415,0.0027590024,0.0003086639,0.0010707679,0.0017205303,0.04976283,0.031963956,0.011899913,0.21140954,0.66768545],"study_design_scores_gemma":[0.00031704252,0.0006664914,0.003758109,0.000141908,0.00008949889,0.00079964637,0.0007517404,0.8665354,0.048884857,0.022119166,0.055832315,0.00010372058],"about_ca_topic_score_codex":0.0033675022,"about_ca_topic_score_gemma":0.0048661125,"teacher_disagreement_score":0.014973057,"about_ca_system_score_codex":0.0008168681,"about_ca_system_score_gemma":0.0019561686,"threshold_uncertainty_score":0.050089836},"labels":[],"label_agreement":null},{"id":"W3203502727","doi":"10.1145/3527546.3527552","title":"A proposed conceptual framework for a representational approach to information retrieval","year":2021,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Ranking (information retrieval); Query expansion; Document retrieval; Similarity (geometry); Relevance (law); Sentence; Natural language; Artificial intelligence; Natural language processing","score_opus":0.038002486838323676,"score_gpt":0.293113914803648,"score_spread":0.2551114279653243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203502727","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012461996,0.0014842161,0.97390896,0.007524277,0.00016866402,0.00021903399,0.00027121455,0.00034577938,0.014831653],"genre_scores_gemma":[0.10135728,0.0026227885,0.8780049,0.0024800242,0.0009413559,0.0015029095,0.0009538402,0.0002301216,0.01190679],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9924758,0.0036620828,0.00071480294,0.0012625182,0.001464065,0.00042062535],"domain_scores_gemma":[0.9940058,0.0026660792,0.000534931,0.0013232948,0.0010821383,0.0003877761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008178432,0.0012373202,0.0011696366,0.00572966,0.0027595875,0.011724974,0.0054630223,0.0044538714,0.011401068],"category_scores_gemma":[0.013328804,0.0009513285,0.002738055,0.0071554505,0.008830754,0.021778038,0.004911011,0.004876734,0.0037784493],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007675441,0.000016754988,0.000060701506,0.00008225188,0.000010751413,0.000028279803,0.0003484046,0.0016863403,0.00024022638,0.9862183,0.0020751155,0.009225172],"study_design_scores_gemma":[0.000020941066,0.00004438664,0.00011157683,0.000096350035,0.000023376626,0.00018374747,0.00033725516,0.01898793,0.0003018735,0.931752,0.04811,0.00003059165],"about_ca_topic_score_codex":0.004946904,"about_ca_topic_score_gemma":0.003068025,"teacher_disagreement_score":0.011724974,"about_ca_system_score_codex":0.004345777,"about_ca_system_score_gemma":0.0039579347,"threshold_uncertainty_score":0.04325217},"labels":[],"label_agreement":null},{"id":"W3203510119","doi":"","title":"AraT5: Text-to-Text Transformers for Arabic Language Understanding and Generation","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Transformer; Computer science; Arabic; Benchmark (surveying); Language model; Natural language processing; Artificial intelligence; Variety (cybernetics); Set (abstract data type); Linguistics; Programming language; Engineering","score_opus":0.15390310461932638,"score_gpt":0.21174942601210256,"score_spread":0.05784632139277618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203510119","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024484627,0.0009582798,0.9237894,0.00090064906,0.00037309548,0.0004678075,0.0028261654,0.03674029,0.009459813],"genre_scores_gemma":[0.4600555,0.00084665284,0.49832454,0.000733323,0.00024310216,0.001086504,0.016188903,0.0029762613,0.019545212],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99883395,0.0004775774,0.00008183182,0.0003299474,0.00018775076,0.00008892279],"domain_scores_gemma":[0.99749136,0.0013937519,0.000108524015,0.00045043643,0.00041497417,0.00014087126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002714478,0.0018987834,0.00069615705,0.0018941684,0.0006515613,0.0019254531,0.002977973,0.0016364721,0.013782259],"category_scores_gemma":[0.0085071875,0.000553218,0.002072618,0.0011173605,0.00076316285,0.0052478625,0.002690881,0.0034399095,0.007484062],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006305039,0.0004035648,0.0024838725,0.000550779,0.00024411913,0.00022690435,0.0004972062,0.16382009,0.009679346,0.02913659,0.050966967,0.74136],"study_design_scores_gemma":[0.000048120233,0.000103126804,0.00032870506,0.000029598245,0.000041654715,0.000085277294,0.000080231635,0.9566805,0.008880519,0.022413736,0.011283141,0.000025360312],"about_ca_topic_score_codex":0.004597061,"about_ca_topic_score_gemma":0.005715901,"teacher_disagreement_score":0.013782259,"about_ca_system_score_codex":0.0013529119,"about_ca_system_score_gemma":0.0015804094,"threshold_uncertainty_score":0.04610628},"labels":[],"label_agreement":null},{"id":"W3203750292","doi":"10.1007/978-3-030-88113-9_46","title":"Cbow Training Time and Accuracy Optimization Using SkipGram","year":2021,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Université du Québec à Montréal","funders":"","keywords":"Training (meteorology); Computer science; Artificial intelligence; Pattern recognition (psychology); Machine learning; Geography; Meteorology","score_opus":0.08994632349629471,"score_gpt":0.3138227792282719,"score_spread":0.22387645573197718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203750292","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053505193,0.0028718149,0.92267054,0.0007961273,0.00042359292,0.00010122241,0.00094959064,0.011941563,0.006740285],"genre_scores_gemma":[0.30998892,0.0007707765,0.65626884,0.00041074608,0.00025500497,0.00024592303,0.004166516,0.0028414375,0.025051793],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859005,0.0003692849,0.00008997123,0.0003594392,0.00035395927,0.00023727423],"domain_scores_gemma":[0.99605227,0.0025969823,0.00010212824,0.00051381305,0.00059333,0.00014140319],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022684378,0.0015599753,0.0016792606,0.0011887545,0.0007973679,0.0014815793,0.0023110833,0.0022830367,0.012177647],"category_scores_gemma":[0.007974502,0.0007607202,0.00072178565,0.0015795233,0.0005501523,0.0031651994,0.0016800538,0.0023075938,0.003964655],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013388956,0.00032326765,0.0015204758,0.00024757703,0.00013151125,0.000117226155,0.00008875054,0.18139759,0.017751625,0.008172026,0.029142035,0.7597691],"study_design_scores_gemma":[0.000025418842,0.00006272849,0.00034402174,0.000012320644,0.000023960114,0.000038382404,0.000020262847,0.9892391,0.006226302,0.0028436687,0.0011550417,0.000008824804],"about_ca_topic_score_codex":0.014175657,"about_ca_topic_score_gemma":0.021774828,"teacher_disagreement_score":0.014175657,"about_ca_system_score_codex":0.0012607253,"about_ca_system_score_gemma":0.0022515333,"threshold_uncertainty_score":0.040738285},"labels":[],"label_agreement":null},{"id":"W3204470331","doi":"10.18653/v1/2022.spnlp-1.7","title":"Predicting Attention Sparsity in Transformers","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Quadratic equation; Softmax function; Transformer; Machine translation; Computation; Encoder; Cluster analysis; Quantization (signal processing); Algorithm; Artificial intelligence; Language model; Theoretical computer science; Machine learning; Mathematics; Artificial neural network","score_opus":0.03754596689072241,"score_gpt":0.25970718152289785,"score_spread":0.22216121463217542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3204470331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17066528,0.00030700813,0.82274836,0.00069401716,0.00003740251,0.000060494178,0.00046755795,0.0025771335,0.0024428319],"genre_scores_gemma":[0.92522407,0.00021987662,0.06957336,0.00020621838,0.0000483014,0.000087602115,0.0007265266,0.0002672821,0.0036466506],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994018,0.00021651629,0.000022282316,0.0001798688,0.00010289522,0.00007656629],"domain_scores_gemma":[0.9941214,0.004460748,0.00026529902,0.00061938586,0.00040464615,0.00012858432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014469712,0.00064111175,0.00070727256,0.0007685837,0.00036231778,0.0011536787,0.0010882522,0.00085376576,0.0032241908],"category_scores_gemma":[0.013963619,0.0004939168,0.0005750138,0.0006988605,0.0009285516,0.0028283657,0.001411328,0.0015724247,0.0007930925],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009064786,0.00012157582,0.00627574,0.00020032194,0.0000764558,0.00018590114,0.0005060531,0.6762547,0.017310424,0.058964264,0.0065297275,0.23266836],"study_design_scores_gemma":[0.000010817035,0.0000213229,0.00037535664,0.0000042994297,0.0000064560554,0.000019781073,0.000015609245,0.9761486,0.0022880847,0.020801445,0.000303651,0.000004508837],"about_ca_topic_score_codex":0.0067256903,"about_ca_topic_score_gemma":0.0109382,"teacher_disagreement_score":0.0067256903,"about_ca_system_score_codex":0.0012330272,"about_ca_system_score_gemma":0.00090242905,"threshold_uncertainty_score":0.013373077},"labels":[],"label_agreement":null},{"id":"W3205028211","doi":"10.1109/icassp43922.2022.9747214","title":"Explainable Fact-Checking Through Question Answering","year":2022,"lang":"en","type":"article","venue":"ICASSP 2022 - 2022 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Question answering; Computer science; Natural language processing; Information retrieval; Programming language; Artificial intelligence","score_opus":0.06308632353957351,"score_gpt":0.31344033151945344,"score_spread":0.25035400797987994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3205028211","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021599438,0.00021371119,0.9604522,0.0022588603,0.000059377497,0.00037822788,0.0007891792,0.012439122,0.001809929],"genre_scores_gemma":[0.399709,0.00034110664,0.5918576,0.0007872769,0.00013984619,0.00037158016,0.0036924675,0.0007952636,0.0023057847],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9863258,0.007618624,0.0009431509,0.0025174748,0.0021075602,0.00048746102],"domain_scores_gemma":[0.91199696,0.06360558,0.004660392,0.015556968,0.0035120696,0.00066803367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01193905,0.0021383977,0.0012344604,0.0024743679,0.0014520226,0.0061729923,0.0045758532,0.0036440666,0.008641964],"category_scores_gemma":[0.06646599,0.0011317666,0.0033194781,0.0016148353,0.0026142672,0.010980691,0.0054714563,0.0037895956,0.0021603876],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001430567,0.00082928257,0.014677586,0.0012579184,0.00069691817,0.0013921817,0.007943026,0.1955552,0.019115679,0.2782657,0.021415217,0.4574206],"study_design_scores_gemma":[0.00006347433,0.00008118396,0.00085090863,0.00011666278,0.00017809217,0.00024817832,0.00031782745,0.8224662,0.0116927335,0.14922963,0.014680665,0.00007450691],"about_ca_topic_score_codex":0.0075557893,"about_ca_topic_score_gemma":0.0055705747,"teacher_disagreement_score":0.01193905,"about_ca_system_score_codex":0.0023440353,"about_ca_system_score_gemma":0.0031184307,"threshold_uncertainty_score":0.06314045},"labels":[],"label_agreement":null},{"id":"W3205430877","doi":"10.1007/978-3-030-88483-3_23","title":"Towards Unifying the Explainability Evaluation Methods for NLP","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Interpretability; Computer science; Artificial intelligence; Phrase; Metric (unit); Set (abstract data type); Machine learning; Feature (linguistics); Focus (optics); Natural language; Deep learning; Key (lock); Natural language processing; Training set","score_opus":0.08998709768257471,"score_gpt":0.3780216545665107,"score_spread":0.288034556883936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3205430877","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002002978,0.0008826519,0.99364626,0.00049281615,0.0000738663,0.00014320489,0.00019407483,0.00089385646,0.0016704184],"genre_scores_gemma":[0.064705685,0.00084362796,0.92960215,0.00022549176,0.00029324892,0.00048666238,0.0011388907,0.00069536874,0.002008914],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97217464,0.013810897,0.0032151134,0.0027292066,0.007384778,0.0006854265],"domain_scores_gemma":[0.8934941,0.075667456,0.003789636,0.013352763,0.012427291,0.0012687797],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03422286,0.0020101871,0.0030041817,0.010424637,0.0012194321,0.010076493,0.0042279386,0.0030767252,0.0057578497],"category_scores_gemma":[0.089282185,0.001313635,0.0039508543,0.0060190465,0.0034428844,0.015564136,0.0075699226,0.0064945556,0.0021062163],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023833149,0.0002464073,0.0041467776,0.0010653747,0.00040334344,0.00009485707,0.0012860153,0.016205674,0.003669543,0.33107036,0.0058301715,0.6357432],"study_design_scores_gemma":[0.00004996031,0.00013361736,0.00256807,0.00051733287,0.0002604426,0.00013667543,0.0004997077,0.33769953,0.004696964,0.6334565,0.019858839,0.0001223061],"about_ca_topic_score_codex":0.0037940843,"about_ca_topic_score_gemma":0.0041956683,"teacher_disagreement_score":0.96577716,"about_ca_system_score_codex":0.0025152673,"about_ca_system_score_gemma":0.0030449326,"threshold_uncertainty_score":0.18098992},"labels":[],"label_agreement":null},{"id":"W3205467272","doi":"10.1145/3534678.3539077","title":"TAG: Toward Accurate Social Media Content Tagging with a Concept Graph","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 28th ACM SIGKDD Conference on Knowledge Discovery and Data Mining","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; University of Alberta","funders":"","keywords":"Computer science; Social media; Conceptualization; Graph; User-generated content; Natural language processing; Artificial intelligence; Information retrieval; World Wide Web; Theoretical computer science","score_opus":0.15618565245343002,"score_gpt":0.2986985138738367,"score_spread":0.1425128614204067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3205467272","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10890982,0.0025597783,0.79545134,0.0014088633,0.0004754209,0.0009778802,0.051520724,0.029966395,0.008729776],"genre_scores_gemma":[0.22059968,0.00067170535,0.6836305,0.0007449963,0.00012258401,0.0006615438,0.08854976,0.0009531633,0.004066143],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99849606,0.00033166338,0.00010518937,0.0006089804,0.00036023185,0.00009790636],"domain_scores_gemma":[0.9973605,0.0012377703,0.00031932123,0.00048662935,0.00048884226,0.0001068275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013869064,0.0015105853,0.00063718425,0.009160057,0.00090662757,0.001462023,0.0020775837,0.002080533,0.0021815223],"category_scores_gemma":[0.0074954336,0.00042840914,0.0013809202,0.0056297313,0.0007951894,0.005537635,0.0021869089,0.0016708871,0.0026676739],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008368722,0.0009564028,0.028063929,0.0022270638,0.0004727228,0.0010280022,0.0018259699,0.04409645,0.03215746,0.05827442,0.1468467,0.6832141],"study_design_scores_gemma":[0.00016135233,0.00023702063,0.00931131,0.00022507318,0.00020217572,0.00084494625,0.0013315233,0.7323746,0.024853904,0.10157284,0.12877432,0.00011090472],"about_ca_topic_score_codex":0.012774567,"about_ca_topic_score_gemma":0.022774756,"teacher_disagreement_score":0.012774567,"about_ca_system_score_codex":0.0015409867,"about_ca_system_score_gemma":0.0017288564,"threshold_uncertainty_score":0.0254004},"labels":[],"label_agreement":null},{"id":"W3205791561","doi":"10.48550/arxiv.2110.07752","title":"Hindsight: Posterior-guided training of retrievers for improved open-ended generation","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Hindsight bias; Computer science; Context (archaeology); Generator (circuit theory); Labrador Retriever; Relevance (law); Information retrieval; Artificial intelligence; Natural language processing; Psychology; Cognitive psychology; Medicine; Power (physics); Geography; Political science","score_opus":0.23432641228749396,"score_gpt":0.23598557641463108,"score_spread":0.001659164127137125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3205791561","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06471445,0.0009137774,0.90726024,0.0007392059,0.00017133758,0.0002228904,0.0004540702,0.020947829,0.0045762802],"genre_scores_gemma":[0.6289619,0.00022598068,0.3523529,0.0013360609,0.00019704351,0.0005455842,0.0031878352,0.002493139,0.010699607],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977883,0.0010581115,0.00011260279,0.00063357374,0.00024020414,0.00016723326],"domain_scores_gemma":[0.9913884,0.0063423947,0.00029551747,0.0010962088,0.0006379255,0.0002394685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048076026,0.0018653248,0.0012577872,0.0007232888,0.0006074205,0.0013347276,0.002960051,0.0026991346,0.00794243],"category_scores_gemma":[0.017885815,0.0009479432,0.0009844261,0.00052369904,0.0012707405,0.0029575005,0.003165159,0.0036975471,0.0036436422],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001612348,0.00068698224,0.0047279987,0.0006401617,0.00023624791,0.0005790161,0.0011949722,0.399337,0.027085472,0.0104627265,0.022452237,0.5309849],"study_design_scores_gemma":[0.00013823502,0.00018540364,0.0004039559,0.00003436262,0.00003465096,0.00012476694,0.00006234484,0.97895056,0.009783867,0.0077566924,0.002500108,0.000025000743],"about_ca_topic_score_codex":0.00311079,"about_ca_topic_score_gemma":0.006122253,"teacher_disagreement_score":0.00794243,"about_ca_system_score_codex":0.0007746724,"about_ca_system_score_gemma":0.001156776,"threshold_uncertainty_score":0.026570082},"labels":[],"label_agreement":null},{"id":"W3205950290","doi":"10.18653/v1/2022.acl-short.17","title":"The Power of Prompt Tuning for Low-Resource Semantic Parsing","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Task (project management); Parsing; Natural language processing; Artificial intelligence; Resource (disambiguation); Meaning (existential); Natural language; Natural language understanding; Language model","score_opus":0.030699814628191713,"score_gpt":0.2701895997179747,"score_spread":0.23948978508978297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3205950290","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19999497,0.0026546845,0.68919694,0.0022562009,0.00084031967,0.00034256623,0.002121984,0.09078996,0.011802395],"genre_scores_gemma":[0.8172649,0.0004580771,0.16903917,0.0013375783,0.00017050975,0.00035076798,0.003742,0.0045940336,0.0030429454],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99769324,0.0009665847,0.00013113853,0.0007594925,0.00026213186,0.00018738955],"domain_scores_gemma":[0.9927248,0.004797875,0.0001925708,0.0015422077,0.00047994562,0.0002625387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004365498,0.0021258339,0.0010041432,0.0007050044,0.00075784844,0.0021770524,0.0022505664,0.0020612786,0.005371679],"category_scores_gemma":[0.02665045,0.0006308519,0.00091517466,0.0009329457,0.0010429636,0.0051503843,0.002476359,0.005072422,0.002821903],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002486061,0.00096874783,0.012522238,0.0009821297,0.0003519212,0.00071575894,0.001081765,0.33915183,0.060189396,0.017204942,0.049382936,0.5149623],"study_design_scores_gemma":[0.0003072646,0.00042964987,0.0019594405,0.000065402324,0.000091118316,0.00023131257,0.00032735438,0.93014073,0.021384737,0.033494227,0.011492671,0.000076227596],"about_ca_topic_score_codex":0.003991793,"about_ca_topic_score_gemma":0.0062196706,"teacher_disagreement_score":0.005371679,"about_ca_system_score_codex":0.0009853616,"about_ca_system_score_gemma":0.002161723,"threshold_uncertainty_score":0.023087263},"labels":[],"label_agreement":null},{"id":"W3206282310","doi":"10.18653/v1/2023.eacl-main.55","title":"What Makes Sentences Semantically Related? A Textual Relatedness Dataset and Empirical Study","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; National Research Council Canada; Institute for Work & Health","funders":"","keywords":"Computer science; Natural language processing; Semantic similarity; Annotation; Artificial intelligence; Automatic summarization; Intuition; Sentence; Similarity (geometry); Information retrieval; Psychology","score_opus":0.09582439625974269,"score_gpt":0.35468479477627646,"score_spread":0.25886039851653375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206282310","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8578566,0.002071379,0.008533373,0.0013136691,0.00013640407,0.0005747357,0.11813592,0.00110753,0.010270405],"genre_scores_gemma":[0.5383066,0.0005146842,0.024589496,0.0004894673,0.0002131829,0.00092210533,0.431438,0.00017934242,0.0033471787],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.99621433,0.00180492,0.00040121746,0.00067608926,0.00074526767,0.00015813309],"domain_scores_gemma":[0.9846021,0.008051129,0.0020038686,0.0019648394,0.0021792443,0.0011987559],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002993324,0.0005800289,0.0003842496,0.003838906,0.0014184349,0.00076477544,0.0010198904,0.0013143816,0.0021252087],"category_scores_gemma":[0.016678696,0.00017629104,0.0005785834,0.003950391,0.00083174935,0.0016190213,0.0013873648,0.0010422623,0.0017177787],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032174801,0.0062465672,0.22438999,0.0063298508,0.00069643767,0.0034828458,0.011366246,0.011083112,0.036616955,0.0148685975,0.44644868,0.23525332],"study_design_scores_gemma":[0.0010619804,0.0021149921,0.6150546,0.0005063589,0.00035156703,0.007218818,0.010108941,0.056645337,0.018001795,0.010537576,0.278093,0.00030499874],"about_ca_topic_score_codex":0.0043528653,"about_ca_topic_score_gemma":0.008346739,"teacher_disagreement_score":0.0043528653,"about_ca_system_score_codex":0.00081552874,"about_ca_system_score_gemma":0.0007169348,"threshold_uncertainty_score":0.015830457},"labels":[],"label_agreement":null},{"id":"W3206407893","doi":"10.21203/rs.3.rs-970738/v1","title":"Machine learning algorithms to identify cluster randomized trials from MEDLINE and EMBASE","year":2021,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Lawson Health Research Institute; McMaster University; London Health Sciences Centre; Ottawa Hospital","funders":"Canadian Institutes of Health Research; Kidney Foundation of Canada; McMaster University; Ottawa Hospital Research Institute","keywords":"MEDLINE; Computer science; Randomized controlled trial; Cluster (spacecraft); Machine learning; Algorithm; Artificial intelligence; Medicine; Internal medicine; Biology","score_opus":0.15620571662061353,"score_gpt":0.4574555004713133,"score_spread":0.3012497838506998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206407893","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2053143,0.033945445,0.7321923,0.003653253,0.00067849027,0.007921765,0.0065976186,0.0059415977,0.0037553182],"genre_scores_gemma":[0.4220604,0.0025900714,0.5635128,0.0007050129,0.0004116412,0.005154518,0.005039302,0.00015780835,0.0003684221],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9232302,0.04655928,0.01726641,0.0057607885,0.0066734985,0.00050992623],"domain_scores_gemma":[0.5010458,0.43705776,0.03158613,0.011828726,0.017563025,0.0009186167],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.113285445,0.001517872,0.0031413154,0.01615664,0.0007600236,0.002981791,0.0021917338,0.0017618412,0.002132504],"category_scores_gemma":[0.41121748,0.00072883937,0.003335144,0.008584666,0.00085454545,0.0026628051,0.001871954,0.0019895534,0.0006545366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049853357,0.00054058287,0.107299134,0.010656836,0.00810528,0.00037834913,0.00065236626,0.12604462,0.0020608401,0.006531164,0.012037684,0.72070783],"study_design_scores_gemma":[0.0024896285,0.0017760828,0.035097122,0.0032459993,0.0041262885,0.00068512635,0.0002266893,0.8718708,0.005967893,0.062582396,0.011708098,0.00022389024],"about_ca_topic_score_codex":0.00215244,"about_ca_topic_score_gemma":0.0036744217,"teacher_disagreement_score":0.8867146,"about_ca_system_score_codex":0.0021062002,"about_ca_system_score_gemma":0.0063616172,"threshold_uncertainty_score":0.5991179},"labels":[],"label_agreement":null},{"id":"W3206487987","doi":"10.18653/v1/2022.acl-long.132","title":"An Empirical Survey of the Effectiveness of Debiasing Techniques for Pre-trained Language Models","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":109,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Institut de Valorisation des Données; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Debiasing; Counterfactual thinking; Computer science; Dropout (neural networks); Machine learning; Artificial intelligence; Psychology; Social psychology","score_opus":0.016971588971216117,"score_gpt":0.2959718698235908,"score_spread":0.2790002808523747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206487987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54458904,0.07442861,0.34154937,0.004115991,0.00069477316,0.00069755263,0.0028202778,0.019535573,0.011568882],"genre_scores_gemma":[0.74113744,0.011220361,0.23528086,0.0009108141,0.00026459314,0.00039621242,0.0066780555,0.0014817319,0.0026299176],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98766905,0.00691467,0.0011312745,0.0019464849,0.001999673,0.00033883596],"domain_scores_gemma":[0.93002963,0.05421005,0.0023080385,0.008962282,0.0038983137,0.00059160794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018784663,0.00238255,0.0012674152,0.0022181694,0.0009484035,0.0016723353,0.0023475036,0.0017597149,0.0015129795],"category_scores_gemma":[0.07769394,0.00081735436,0.0014086725,0.0017578237,0.0014085848,0.0042899144,0.0023369084,0.0034136565,0.0014253597],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010623366,0.00043605422,0.035093285,0.0029342924,0.0012368779,0.00013643787,0.0009953333,0.080453575,0.011769029,0.0031218417,0.012766318,0.8499946],"study_design_scores_gemma":[0.0002920426,0.0028679208,0.03830665,0.0021141618,0.0011204183,0.0011969733,0.0015448155,0.8164682,0.08681882,0.011594843,0.037381753,0.0002935065],"about_ca_topic_score_codex":0.0048113232,"about_ca_topic_score_gemma":0.0070907557,"teacher_disagreement_score":0.018784663,"about_ca_system_score_codex":0.0013260399,"about_ca_system_score_gemma":0.0016137596,"threshold_uncertainty_score":0.099343956},"labels":[],"label_agreement":null},{"id":"W3206557162","doi":"","title":"PRIMER: Pyramid-based Masked Sentence Pre-training for Multi-document Summarization.","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Pyramid (geometry); Artificial intelligence; Salient; Natural language processing; Code (set theory); Multi-document summarization; Focus (optics); Information retrieval; Set (abstract data type)","score_opus":0.13837903921162628,"score_gpt":0.23164147337613547,"score_spread":0.09326243416450919,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206557162","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017357014,0.0017883072,0.93251586,0.00040083102,0.0003259431,0.00039420466,0.0026871066,0.041795623,0.0027351344],"genre_scores_gemma":[0.21727522,0.0008364309,0.74002236,0.0008832861,0.00026453592,0.001164952,0.021983553,0.0022365127,0.015333086],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994468,0.0001567051,0.000034710538,0.00022508018,0.00008156864,0.000055127824],"domain_scores_gemma":[0.9989042,0.00052555854,0.00007991727,0.0001913105,0.0002291065,0.00006987767],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011010184,0.0017484489,0.0007639407,0.00094562356,0.00042884817,0.0007704143,0.002345818,0.0014308344,0.0061653005],"category_scores_gemma":[0.0037228935,0.0005963032,0.001030138,0.00083403615,0.0003914202,0.0021064107,0.00124432,0.0023417873,0.004936877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043055118,0.000351758,0.001273198,0.0007757626,0.00024103733,0.00024809505,0.00040206808,0.09732729,0.03882291,0.0041741496,0.064053215,0.79190004],"study_design_scores_gemma":[0.00007045023,0.00031436203,0.00072842836,0.000042502994,0.00008450794,0.00012814722,0.0001006721,0.9504794,0.026492257,0.0070010633,0.0145240575,0.000034245615],"about_ca_topic_score_codex":0.003968798,"about_ca_topic_score_gemma":0.010276,"teacher_disagreement_score":0.0061653005,"about_ca_system_score_codex":0.0008100032,"about_ca_system_score_gemma":0.001209827,"threshold_uncertainty_score":0.020624995},"labels":[],"label_agreement":null},{"id":"W3206907172","doi":"10.18653/v1/2022.acl-long.125","title":"Composable Sparse Fine-Tuning for Cross-Lingual Transfer","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Language model; Fine-tuning; Modular design; Overfitting; Artificial intelligence; Task (project management); Inference; Adapter (computing); Programming language; Computer hardware","score_opus":0.01407345951011647,"score_gpt":0.2517387941339686,"score_spread":0.23766533462385214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206907172","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055678677,0.0009672241,0.9153187,0.000518082,0.00020521265,0.00018304808,0.0005402852,0.02180029,0.0047883075],"genre_scores_gemma":[0.69945866,0.00037068388,0.28564858,0.0010987741,0.00013087397,0.0005800321,0.0025154226,0.0027212629,0.0074756127],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988733,0.00028304438,0.0000611115,0.0004353212,0.00019931531,0.00014792263],"domain_scores_gemma":[0.9981269,0.0008586173,0.0000905683,0.00061313587,0.0002212138,0.000089633424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019615784,0.0019517981,0.0012908299,0.00064595946,0.00064788834,0.0011912555,0.0033367155,0.0017360491,0.0063976557],"category_scores_gemma":[0.008000606,0.0009639036,0.0013133535,0.0006471336,0.0011250579,0.003107284,0.0031318972,0.0040328265,0.0032931971],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004353038,0.0003956809,0.0030861967,0.00029473574,0.00035280024,0.000282644,0.00036114326,0.52708733,0.021690212,0.009308142,0.014265555,0.4224402],"study_design_scores_gemma":[0.000023419474,0.000046468125,0.0002545232,0.000012993025,0.00002329089,0.000042216405,0.000030256224,0.98472637,0.0037753165,0.009451046,0.0015981571,0.000016037759],"about_ca_topic_score_codex":0.008811509,"about_ca_topic_score_gemma":0.015870387,"teacher_disagreement_score":0.008811509,"about_ca_system_score_codex":0.001191861,"about_ca_system_score_gemma":0.0014995128,"threshold_uncertainty_score":0.0214023},"labels":[],"label_agreement":null},{"id":"W3206996280","doi":"10.18653/v1/2022.findings-naacl.184","title":"CCQA: A New Web-Scale Question Answering Dataset for Model Pre-Training","year":2022,"lang":"en","type":"article","venue":"Findings of the Association for Computational Linguistics: NAACL 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Question answering; Computer science; Open domain; Task (project management); Language model; Domain (mathematical analysis); Artificial intelligence; Natural language processing; Scale (ratio); Training set; Information retrieval; Resource (disambiguation); Natural language; Natural language understanding","score_opus":0.024946483409200453,"score_gpt":0.2847366473702726,"score_spread":0.25979016396107213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206996280","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12543306,0.0044658626,0.12004773,0.0033744194,0.0012015299,0.0028411355,0.6442641,0.08274676,0.015625397],"genre_scores_gemma":[0.0571815,0.00034489608,0.073013596,0.0008548268,0.00016924723,0.0017994597,0.86189425,0.0010798576,0.0036624174],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962774,0.001265941,0.00036322634,0.0011188239,0.0007297485,0.00024496793],"domain_scores_gemma":[0.9919264,0.0030654683,0.00037140463,0.0020403767,0.0019216799,0.00067472784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032401495,0.0027537001,0.0013396178,0.0042916546,0.001685424,0.0019022243,0.0041624773,0.0035417746,0.0070302594],"category_scores_gemma":[0.014567032,0.0007344283,0.0021314938,0.0034177843,0.0009826336,0.0036684566,0.0038213255,0.0036911604,0.010096237],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067070464,0.001996464,0.011246784,0.0027336332,0.00033279418,0.0005369476,0.0010447161,0.012562797,0.012868493,0.0050311196,0.8091158,0.14185971],"study_design_scores_gemma":[0.0012488556,0.0010007857,0.034746327,0.00045643596,0.00034814785,0.0016330238,0.0014981083,0.28053483,0.030064585,0.021643553,0.6264046,0.0004207354],"about_ca_topic_score_codex":0.021281388,"about_ca_topic_score_gemma":0.040030234,"teacher_disagreement_score":0.021281388,"about_ca_system_score_codex":0.0019118458,"about_ca_system_score_gemma":0.003238328,"threshold_uncertainty_score":0.042315066},"labels":[],"label_agreement":null},{"id":"W3207058314","doi":"10.1609/aaai.v36i10.21308","title":"Language Modelling via Learning to Rank","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Vector Institute","funders":"Natural Sciences and Engineering Research Council of Canada; Vector Institute","keywords":"Perplexity; Computer science; Language model; Transformer; Natural language processing; Artificial intelligence; Probabilistic logic; Entropy (arrow of time); Machine learning; Rank (graph theory); Ranking (information retrieval); Mathematics","score_opus":0.07286834451957244,"score_gpt":0.28741439635267724,"score_spread":0.2145460518331048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207058314","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010537541,0.0004543177,0.97929907,0.0006839253,0.00008723248,0.000057340316,0.0005436781,0.00526363,0.0030733286],"genre_scores_gemma":[0.5003977,0.00069935026,0.4775813,0.0007265804,0.00033677623,0.0003972575,0.0039788596,0.000930877,0.014951234],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99811816,0.00091365754,0.000082174935,0.00043769053,0.00030636313,0.0001418636],"domain_scores_gemma":[0.99527234,0.0030164823,0.00026062346,0.00070657424,0.0006019359,0.0001420862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019977742,0.0013080755,0.00088231545,0.0013111987,0.0006316051,0.0022150867,0.0019199764,0.0014070261,0.0056891777],"category_scores_gemma":[0.009500859,0.00048622064,0.0011224215,0.0009962915,0.00077642064,0.0030659896,0.0018350378,0.0025154396,0.005738337],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023750037,0.00016835357,0.0016680283,0.0002890913,0.00013990144,0.00017526273,0.000270066,0.46762595,0.004981708,0.045905877,0.018819053,0.45971927],"study_design_scores_gemma":[0.000013242231,0.00003437788,0.00010705865,0.000012718179,0.000010389496,0.000034491244,0.000019893509,0.95900536,0.0017184786,0.0369519,0.0020763117,0.000015744055],"about_ca_topic_score_codex":0.0049154665,"about_ca_topic_score_gemma":0.008937541,"teacher_disagreement_score":0.0056891777,"about_ca_system_score_codex":0.0010604989,"about_ca_system_score_gemma":0.001312722,"threshold_uncertainty_score":0.01903218},"labels":[],"label_agreement":null},{"id":"W3207166518","doi":"10.18653/v1/2022.naacl-main.341","title":"Symbolic Knowledge Distillation: from General Language Models to Commonsense Models","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":150,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Naval Information Warfare Center Pacific; Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; Allen Institute for Artificial Intelligence","keywords":"Computer science; Computational linguistics; Language model; Artificial intelligence; Natural language processing; Cognitive science; Linguistics; Philosophy; Psychology","score_opus":0.029680022821880935,"score_gpt":0.2597489079193148,"score_spread":0.2300688850974339,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207166518","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014780297,0.0013307735,0.9710153,0.002328981,0.00016029115,0.00006852211,0.0008793494,0.0024054581,0.0070310896],"genre_scores_gemma":[0.53968054,0.0016690561,0.44529176,0.0007072332,0.0003063484,0.00035510916,0.0031192873,0.0009610816,0.0079094395],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987614,0.00058171066,0.00007014883,0.0002221642,0.0002592082,0.00010536694],"domain_scores_gemma":[0.9930997,0.005292176,0.00018445744,0.00080237375,0.0004468377,0.00017448065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019989116,0.0010466464,0.0014281635,0.0016494171,0.0011573496,0.0035592334,0.003037818,0.0014433792,0.009483786],"category_scores_gemma":[0.013583582,0.00089824916,0.0016802832,0.0021421227,0.0021915003,0.009962929,0.004991059,0.0042710076,0.0020930504],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004615095,0.00018673675,0.00073249656,0.0004113901,0.00015391633,0.0002019563,0.0005431865,0.30580866,0.0009029977,0.3870051,0.014606584,0.28898534],"study_design_scores_gemma":[0.000023922155,0.000012577101,0.000040158964,0.000033524484,0.000015487103,0.00001635209,0.000041746567,0.6219463,0.00049055053,0.3751429,0.0022235098,0.00001292855],"about_ca_topic_score_codex":0.0083525395,"about_ca_topic_score_gemma":0.015389774,"teacher_disagreement_score":0.009483786,"about_ca_system_score_codex":0.0013626505,"about_ca_system_score_gemma":0.0019574556,"threshold_uncertainty_score":0.03172642},"labels":[],"label_agreement":null},{"id":"W3207553988","doi":"10.18653/v1/2022.acl-long.225","title":"Generated Knowledge Prompting for Commonsense Reasoning","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":176,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Naval Information Warfare Center Pacific; Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; Allen Institute for Artificial Intelligence","keywords":"Commonsense reasoning; Commonsense knowledge; Cognitive science; Computer science; Computational linguistics; Artificial intelligence; Natural language processing; Linguistics; Epistemology; Philosophy; Psychology; Knowledge representation and reasoning","score_opus":0.013502953991089174,"score_gpt":0.247305549832438,"score_spread":0.23380259584134883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207553988","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040156186,0.0014726508,0.86694425,0.0027764111,0.0011413194,0.0006738616,0.0045923423,0.04940558,0.032837413],"genre_scores_gemma":[0.5869941,0.00056540343,0.39065796,0.0007994191,0.0002815887,0.0003631262,0.0076107876,0.0013958289,0.011331772],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99812645,0.0006853061,0.00011395851,0.0004772819,0.0004504469,0.00014648464],"domain_scores_gemma":[0.99355066,0.0041596624,0.00020015583,0.0011283742,0.00070328446,0.00025782146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021819947,0.00089420355,0.0007699031,0.0015924734,0.0010018417,0.0021738228,0.002112424,0.0014950986,0.023758706],"category_scores_gemma":[0.012157863,0.00038557462,0.0007501617,0.0009789519,0.0007863146,0.0050906427,0.004671333,0.0019968173,0.004629149],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026232167,0.00068331184,0.0018106695,0.0011129411,0.00011392527,0.0015123244,0.0017767685,0.01838575,0.02118055,0.14072956,0.10060039,0.7094706],"study_design_scores_gemma":[0.0004277672,0.00021431329,0.00083962915,0.0002640275,0.00017045255,0.00043286217,0.0008164878,0.44871777,0.04578713,0.37406912,0.1281798,0.00008062512],"about_ca_topic_score_codex":0.0016277272,"about_ca_topic_score_gemma":0.0032753572,"teacher_disagreement_score":0.023758706,"about_ca_system_score_codex":0.001075509,"about_ca_system_score_gemma":0.0012868728,"threshold_uncertainty_score":0.07948083},"labels":[],"label_agreement":null},{"id":"W3207783082","doi":"","title":"BI-RADS BERT & Using Section Tokenization to Understand Radiology Reports.","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Health Sciences Centre; Sunnybrook Health Science Centre; University of Toronto","funders":"","keywords":"Lexical analysis; Computer science; Artificial intelligence; Lexicon; Natural language processing; Preprocessor; Breast imaging; Section (typography); Radiology; Breast cancer; Mammography; Medicine; Cancer","score_opus":0.11669552856201919,"score_gpt":0.20118952858465539,"score_spread":0.0844940000226362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207783082","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.099545635,0.0013339741,0.86919457,0.0011510843,0.00043825712,0.00031083782,0.0064069508,0.014463424,0.0071551907],"genre_scores_gemma":[0.63431096,0.001037357,0.3283456,0.0003300045,0.00025329436,0.00031360608,0.020885529,0.001061892,0.013461756],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993254,0.00023475646,0.000070046,0.00020179455,0.00011728966,0.00005066778],"domain_scores_gemma":[0.99790347,0.00097106706,0.00031096497,0.0004386635,0.00031776057,0.000058159683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001346427,0.0008208808,0.00024312666,0.0012288264,0.0002373674,0.0010311061,0.00058231555,0.0005222472,0.0028186953],"category_scores_gemma":[0.004843607,0.00034489247,0.000875766,0.0007073406,0.00031996134,0.0028930316,0.0009201461,0.00091841584,0.0034785068],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009253332,0.00029174492,0.040349755,0.0008330695,0.00025075887,0.0007119549,0.0013275115,0.095576346,0.047970716,0.03108102,0.03894969,0.74173194],"study_design_scores_gemma":[0.000029030907,0.00027734807,0.018495498,0.00013120813,0.00013880125,0.001024487,0.0003755736,0.8364249,0.037190143,0.029489642,0.07633642,0.000087008586],"about_ca_topic_score_codex":0.005270118,"about_ca_topic_score_gemma":0.0074141365,"teacher_disagreement_score":0.005270118,"about_ca_system_score_codex":0.000729126,"about_ca_system_score_gemma":0.000983939,"threshold_uncertainty_score":0.010478914},"labels":[],"label_agreement":null},{"id":"W3207937903","doi":"10.1162/tacl_a_00416","title":"MasakhaNER: Named Entity Recognition for African Languages","year":2021,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":235,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Google (Canada)","funders":"","keywords":"Computer science; Named-entity recognition; Variety (cybernetics); Natural language processing; Artificial intelligence; Representation (politics); Code (set theory); Quality (philosophy); Data science; Information retrieval; Programming language; Task (project management); Political science","score_opus":0.024774853785437097,"score_gpt":0.27454927867574097,"score_spread":0.24977442489030388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207937903","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12734428,0.0051169107,0.100572236,0.005284137,0.0019121146,0.0018048099,0.5555237,0.18422633,0.018215438],"genre_scores_gemma":[0.1003961,0.0009504109,0.10636958,0.00058472034,0.00022075714,0.0013371044,0.7819177,0.0018415326,0.0063821203],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9968129,0.0011336047,0.0003197939,0.0008801511,0.00054324174,0.00031032934],"domain_scores_gemma":[0.9954313,0.0015477785,0.00036770187,0.001680911,0.00061035616,0.0003618619],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004316772,0.0022081567,0.0013965566,0.0047165514,0.0021673723,0.0027111422,0.0028022912,0.0019764446,0.012045322],"category_scores_gemma":[0.0110689765,0.0007505238,0.001626501,0.003952262,0.00084892125,0.0073099267,0.0051943846,0.0030490248,0.015280195],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017969569,0.00081522035,0.013343049,0.0021741376,0.00044588756,0.000940068,0.00095022394,0.014205088,0.012753737,0.010503031,0.72939986,0.2126728],"study_design_scores_gemma":[0.00096690655,0.0006499147,0.030581031,0.00070554705,0.00040052994,0.0021771751,0.0019073321,0.26418003,0.054028288,0.02406138,0.6198926,0.00044921867],"about_ca_topic_score_codex":0.0118306065,"about_ca_topic_score_gemma":0.013772555,"teacher_disagreement_score":0.012045322,"about_ca_system_score_codex":0.0011111832,"about_ca_system_score_gemma":0.0025695954,"threshold_uncertainty_score":0.0402956},"labels":[],"label_agreement":null},{"id":"W3208114106","doi":"10.3389/fcomp.2021.674333","title":"Design and Analysis of a Collaborative Story Generation Game for Social Robots","year":2021,"lang":"en","type":"article","venue":"Frontiers in Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Storytelling; Computer science; Entertainment; Robot; Quality (philosophy); Baseline (sea); Artificial intelligence; Multimedia; Human–computer interaction; Art; Narrative; Visual arts","score_opus":0.037368744310425694,"score_gpt":0.2728842541574876,"score_spread":0.23551550984706193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208114106","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44918734,0.0004765314,0.5193775,0.0009065634,0.00013586377,0.007229516,0.0010610208,0.0039884686,0.017637223],"genre_scores_gemma":[0.5881866,0.00014061769,0.4010972,0.00013635648,0.0000143966,0.0025949741,0.0009792414,0.00026054622,0.0065900343],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982919,0.00083869725,0.00009648914,0.00025639316,0.00037627848,0.00014018604],"domain_scores_gemma":[0.9965725,0.0022202544,0.00020413945,0.00018589023,0.00038214106,0.00043511513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018305646,0.0010745232,0.00055870996,0.00073812995,0.0007952777,0.0016266124,0.0018587905,0.001256159,0.007165746],"category_scores_gemma":[0.007165407,0.0005137341,0.0006699028,0.0002454172,0.0007926694,0.0015031992,0.001282361,0.0011762539,0.0011844396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0057169925,0.0072072577,0.017675221,0.0041333684,0.00045421417,0.0047441167,0.010146837,0.45362115,0.13827188,0.06061022,0.014712668,0.28270614],"study_design_scores_gemma":[0.0006347194,0.0025183056,0.0036716075,0.000067708825,0.00007573055,0.0003847041,0.0014430339,0.94571984,0.015440836,0.009068533,0.020883583,0.00009152707],"about_ca_topic_score_codex":0.0026667921,"about_ca_topic_score_gemma":0.004561638,"teacher_disagreement_score":0.007165746,"about_ca_system_score_codex":0.0012368972,"about_ca_system_score_gemma":0.0010087024,"threshold_uncertainty_score":0.023971856},"labels":[],"label_agreement":null},{"id":"W3208143504","doi":"10.5281/zenodo.4073433","title":"Preprocessing scripts and data for study: Identifying high-confidence capture Hi-C interactions using CHiCANE","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Cancer Research","funders":"","keywords":"Scripting language; Preprocessor; Computer science; Operating system; Programming language","score_opus":0.160977707385079,"score_gpt":0.3419859801446908,"score_spread":0.1810082727596118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208143504","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013324919,0.0001243491,0.001509354,0.000052070027,0.00003161594,0.00010238691,0.9916889,0.0044368496,0.0007221377],"genre_scores_gemma":[0.0012248089,0.000040492283,0.0025308381,0.000044026747,0.0000062221893,0.00040024298,0.995,0.00025676354,0.00049654365],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990841,0.00012877153,0.00010547298,0.00041068066,0.00015628309,0.00011468183],"domain_scores_gemma":[0.9983041,0.0006045254,0.00013137131,0.00046995838,0.00033658138,0.00015342906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013449205,0.0028929545,0.0012359045,0.0027889062,0.00091513444,0.0016187254,0.0020758673,0.0015635503,0.030544043],"category_scores_gemma":[0.0040552123,0.0006469366,0.0013576839,0.0032599308,0.00042180644,0.00061302475,0.0010827014,0.0014829428,0.038359087],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042147873,0.00013315518,0.004335882,0.0017381398,0.000106169166,0.00011453937,0.00008391051,0.0013624079,0.0041001816,0.00079533947,0.97573984,0.0110690445],"study_design_scores_gemma":[0.00083478924,0.00018238873,0.01919672,0.00031363647,0.00023019032,0.00044004526,0.00022023603,0.0081079025,0.012353952,0.004623439,0.9533881,0.000108563036],"about_ca_topic_score_codex":0.011743771,"about_ca_topic_score_gemma":0.023113752,"teacher_disagreement_score":0.030544043,"about_ca_system_score_codex":0.0010710066,"about_ca_system_score_gemma":0.0023218864,"threshold_uncertainty_score":0.102180004},"labels":[],"label_agreement":null},{"id":"W3208809827","doi":"10.1145/3459637.3481910","title":"Dual Learning for Query Generation and Query Selection in Query Feeds Recommendation","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Natural Science Foundation of China","keywords":"Computer science; Information retrieval; Selection (genetic algorithm); Generator (circuit theory); Web query classification; Readability; Filter (signal processing); Query expansion; Web search query; Sargable; Query optimization; Query language; Data mining; Search engine; Machine learning","score_opus":0.03996731213320634,"score_gpt":0.2708048305776304,"score_spread":0.23083751844442404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208809827","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035861906,0.0008277236,0.9578473,0.000529486,0.00006537034,0.0003184874,0.0003559219,0.0029894754,0.0012042674],"genre_scores_gemma":[0.5406891,0.00055334513,0.44810754,0.0006038435,0.00033145075,0.0007493109,0.0027662928,0.00032947914,0.0058696945],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955343,0.0019677854,0.00033050493,0.0010362688,0.0008294144,0.0003017063],"domain_scores_gemma":[0.9890175,0.0074282805,0.0005636748,0.0012169223,0.0014211939,0.00035248575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059686094,0.0015050389,0.0020780496,0.0027818931,0.00084825495,0.0016967715,0.0028935669,0.0021286376,0.00233951],"category_scores_gemma":[0.016537294,0.00079717796,0.0013269009,0.0023139447,0.0011821616,0.0038272352,0.002278889,0.002246018,0.0017686657],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016219046,0.001563513,0.009774832,0.0005039793,0.00021301518,0.00030171912,0.0005015225,0.14715284,0.024874909,0.00808779,0.01179599,0.793608],"study_design_scores_gemma":[0.00006561578,0.00015538519,0.0005281971,0.000007833638,0.000030055762,0.00008714491,0.00003366773,0.9895705,0.00473044,0.0035649615,0.0012035079,0.00002271393],"about_ca_topic_score_codex":0.0057828557,"about_ca_topic_score_gemma":0.0071859797,"teacher_disagreement_score":0.0059686094,"about_ca_system_score_codex":0.0012002067,"about_ca_system_score_gemma":0.0016121238,"threshold_uncertainty_score":0.031565428},"labels":[],"label_agreement":null},{"id":"W3208816465","doi":"10.5281/zenodo.5504194","title":"ExploreASL/ExploreASL: ExploreASL v1.8.0","year":2021,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science","score_opus":0.07513490076175912,"score_gpt":0.2492624769311891,"score_spread":0.17412757616942998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208816465","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0007513457,0.0003080667,0.050954983,0.0006259283,0.00033884266,0.00025403206,0.059507214,0.8447351,0.04252444],"genre_scores_gemma":[0.008580919,0.0003911591,0.043771915,0.00090617535,0.00013734617,0.0010802762,0.13101973,0.7591201,0.05499241],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99704844,0.0005329915,0.0002667532,0.00056345714,0.0012303522,0.00035817464],"domain_scores_gemma":[0.99585295,0.0010198767,0.00020851636,0.0010708857,0.001500534,0.0003472637],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0034257586,0.0038117678,0.0015776707,0.0030575653,0.0013148317,0.0072255805,0.0063276356,0.0025612887,0.40757135],"category_scores_gemma":[0.013148567,0.0031696304,0.003069033,0.0024607456,0.0013033989,0.00797214,0.0067302943,0.004793199,0.45935264],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029542323,0.00003903127,0.00036724313,0.0007263673,0.00004108723,0.00010721093,0.00024777258,0.0006222065,0.0021292646,0.006632918,0.953996,0.0347955],"study_design_scores_gemma":[0.0001434913,0.000038775157,0.00044391942,0.00020075782,0.00002759858,0.00015940776,0.00010765567,0.0024975066,0.0057437504,0.009469816,0.9810524,0.000114804745],"about_ca_topic_score_codex":0.0057013584,"about_ca_topic_score_gemma":0.0052409,"teacher_disagreement_score":0.40757135,"about_ca_system_score_codex":0.0017846243,"about_ca_system_score_gemma":0.0027984257,"threshold_uncertainty_score":0.84502757},"labels":[],"label_agreement":null},{"id":"W3208821253","doi":"10.2200/s01123ed1v01y202108hlt053","title":"Pretrained Transformers for Text Ranking: BERT and Beyond","year":2021,"lang":"en","type":"article","venue":"Synthesis lectures on human language technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":124,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Transformer; Computer science; Artificial intelligence; Natural language processing; Encoder; Ranking (information retrieval); Question answering; Information retrieval; Artificial neural network; Engineering","score_opus":0.020761885001247454,"score_gpt":0.2760993430802039,"score_spread":0.25533745807895647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208821253","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0119023565,0.002534807,0.95881015,0.001192462,0.00058442506,0.0001313437,0.0022852668,0.016801743,0.0057573603],"genre_scores_gemma":[0.3954889,0.003056148,0.549694,0.0009845981,0.001948993,0.00038652823,0.017045379,0.0037563988,0.02763904],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99736625,0.000870541,0.0001983626,0.000641576,0.0006601646,0.00026328544],"domain_scores_gemma":[0.9911221,0.0043860096,0.00022335362,0.0025833051,0.0013881832,0.0002970955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003520818,0.0019293224,0.002056973,0.0026898037,0.0008967787,0.003632459,0.002754181,0.0018857123,0.019479072],"category_scores_gemma":[0.017923076,0.000881967,0.0012247313,0.0026592049,0.00094293756,0.010027232,0.0024417252,0.0039595123,0.012888979],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007045155,0.00021981579,0.00096942234,0.0003401917,0.000087781824,0.00007280132,0.00010131357,0.024661427,0.006132819,0.045282673,0.06565246,0.85577476],"study_design_scores_gemma":[0.00011918843,0.00021951752,0.00070941215,0.00010407307,0.000080157035,0.00017041917,0.00010911133,0.78150195,0.011842371,0.18222779,0.02286509,0.000050994033],"about_ca_topic_score_codex":0.005368648,"about_ca_topic_score_gemma":0.0102284085,"teacher_disagreement_score":0.019479072,"about_ca_system_score_codex":0.0013151573,"about_ca_system_score_gemma":0.0025923857,"threshold_uncertainty_score":0.06516397},"labels":[],"label_agreement":null},{"id":"W3209039755","doi":"10.18653/v1/2021.emnlp-main.821","title":"IndoNLI: A Natural Language Inference Dataset for Indonesian","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Universitas Indonesia; York University; Samsung; National Science Foundation","keywords":"Annotation; Indonesian; Sentence; Test set; Set (abstract data type); Inference; Data set","score_opus":0.06648941278896932,"score_gpt":0.43358885394360996,"score_spread":0.36709944115464066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209039755","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057119317,0.002061597,0.021396972,0.00120106,0.00040400636,0.0008615474,0.87964547,0.014065501,0.023244623],"genre_scores_gemma":[0.038948115,0.00025855814,0.023245199,0.00032630155,0.000050104096,0.0008208487,0.93276006,0.00038241036,0.0032084424],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981651,0.00048574424,0.00027141912,0.0006153373,0.0003463194,0.00011615258],"domain_scores_gemma":[0.9973979,0.0010062184,0.00027621837,0.00061351416,0.00050937984,0.00019681036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013626796,0.0014640612,0.00083907286,0.0027366038,0.0012735211,0.0012346931,0.0018473407,0.0016589052,0.010813938],"category_scores_gemma":[0.005044593,0.00048020587,0.0007408656,0.002232438,0.0006360148,0.0021403967,0.00179399,0.002141649,0.011376153],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007707928,0.00096703053,0.017499242,0.0052464013,0.00024111499,0.0021567543,0.0022569357,0.006330294,0.01786698,0.006519645,0.79559696,0.1445479],"study_design_scores_gemma":[0.00034974242,0.0002511561,0.08519321,0.0006117635,0.00019634474,0.0026465005,0.0026624908,0.050623953,0.022610486,0.009092948,0.8254596,0.0003017908],"about_ca_topic_score_codex":0.013645105,"about_ca_topic_score_gemma":0.032189466,"teacher_disagreement_score":0.013645105,"about_ca_system_score_codex":0.0017446107,"about_ca_system_score_gemma":0.0022054038,"threshold_uncertainty_score":0.036176264},"labels":[],"label_agreement":null},{"id":"W3209235691","doi":"10.2196/32698","title":"A BERT-Based Generation Model to Transform Medical Texts to SQL Queries for Electronic Medical Records: Model Development and Validation","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Natural Science Foundation of China; Baidu","keywords":"Computer science; SQL; Information retrieval; Stored procedure; Query by Example; Programming language; Natural language processing; Artificial intelligence; Database; Data mining; Search engine; Web search query","score_opus":0.03404106744262877,"score_gpt":0.30363756947698667,"score_spread":0.2695965020343579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209235691","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20554413,0.002047195,0.773559,0.001974038,0.00029148237,0.00073982007,0.0020677052,0.0078409165,0.0059357216],"genre_scores_gemma":[0.8015644,0.00087928097,0.18447022,0.00063167646,0.00010554949,0.0009822028,0.004162268,0.0001846965,0.0070196246],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99959177,0.00010273199,0.00003090878,0.00015098206,0.000068884685,0.000054805703],"domain_scores_gemma":[0.99844235,0.0009787511,0.000081009435,0.00008121505,0.0003641733,0.000052374202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014608226,0.0011621636,0.0005978613,0.0010810054,0.00037329653,0.00087715936,0.0016281488,0.0012453314,0.003401586],"category_scores_gemma":[0.0041219597,0.00044243457,0.0011782225,0.0005909979,0.0004669226,0.001421841,0.00078809354,0.0018695901,0.0010228072],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032712435,0.00029341842,0.0051179696,0.0001480668,0.00012875104,0.0002341429,0.00013914047,0.76528186,0.0033577716,0.003303209,0.0046175243,0.21705095],"study_design_scores_gemma":[0.000008281144,0.000024121911,0.00015588608,0.0000053016956,0.000012197442,0.000014389297,0.0000066029206,0.9984621,0.00051142625,0.00065042474,0.0001460959,0.0000033263527],"about_ca_topic_score_codex":0.023507489,"about_ca_topic_score_gemma":0.02192361,"teacher_disagreement_score":0.023507489,"about_ca_system_score_codex":0.0016484554,"about_ca_system_score_gemma":0.0020484407,"threshold_uncertainty_score":0.046741307},"labels":[],"label_agreement":null},{"id":"W3209857478","doi":"10.4018/ijirr.289950","title":"An End-to-End Efficient Lucene-Based Framework of Document/Information Retrieval","year":2021,"lang":"en","type":"article","venue":"International Journal of Information Retrieval Research","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Université du Québec à Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Word (group theory); Automatic summarization; Context (archaeology); Set (abstract data type); Document retrieval; Natural language processing","score_opus":0.03288790709535719,"score_gpt":0.37630327528449964,"score_spread":0.3434153681891424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209857478","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031445858,0.00025179822,0.98624057,0.00008684984,0.000025862622,0.00011490449,0.00034417244,0.008506101,0.0012851694],"genre_scores_gemma":[0.073551744,0.00048480232,0.9144794,0.00016273493,0.00007830269,0.00034002613,0.0030215613,0.0008204538,0.00706097],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987973,0.00025240894,0.000119391,0.00025181385,0.00048794236,0.00009113529],"domain_scores_gemma":[0.9986707,0.00031098185,0.000085635984,0.00045843696,0.00041608294,0.00005826082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014716821,0.0007814964,0.0009537662,0.0017939003,0.00086642674,0.0028261985,0.001640215,0.0011574152,0.0050510685],"category_scores_gemma":[0.0038598601,0.00039815973,0.0009984776,0.0015807616,0.00069801643,0.0043446072,0.0020247647,0.0013833967,0.006852755],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010223218,0.00046604377,0.0014417557,0.0012002412,0.00027490867,0.00080141245,0.001365187,0.04758091,0.17793483,0.10080604,0.0361764,0.63092995],"study_design_scores_gemma":[0.000094939234,0.00047316396,0.0010265158,0.00011806277,0.0001366064,0.00096407917,0.0002657596,0.6585739,0.17709972,0.053447206,0.10759803,0.0002019655],"about_ca_topic_score_codex":0.0020311312,"about_ca_topic_score_gemma":0.0027797166,"teacher_disagreement_score":0.0050510685,"about_ca_system_score_codex":0.0007535575,"about_ca_system_score_gemma":0.0008467468,"threshold_uncertainty_score":0.01689756},"labels":[],"label_agreement":null},{"id":"W3210290009","doi":"10.1145/3459637.3482135","title":"Location-Aware Named Entity Disambiguation","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Huawei Technologies","keywords":"Computer science; Entity linking; Inference; Information retrieval; Artificial intelligence; Natural language processing; Dimension (graph theory); Embedding; Baseline (sea); Named-entity recognition; Natural language; Named entity; Knowledge base; Task (project management)","score_opus":0.023266838054333543,"score_gpt":0.25560297465951964,"score_spread":0.2323361366051861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210290009","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01854807,0.0027537665,0.95501584,0.0006482269,0.0003450856,0.00015488769,0.0039383597,0.013453746,0.005141905],"genre_scores_gemma":[0.2712269,0.002119745,0.6997182,0.000558927,0.00040025252,0.00013006521,0.016670307,0.0007661392,0.008409355],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997769,0.0005397201,0.00025123166,0.0008221233,0.00047329784,0.00014468908],"domain_scores_gemma":[0.99508375,0.0016115281,0.00053442724,0.0017447604,0.0008830766,0.00014251031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024477362,0.0010456388,0.0014089771,0.0050259405,0.0012364214,0.0023278366,0.0020619624,0.0014760133,0.0028056165],"category_scores_gemma":[0.0069708787,0.0005135177,0.0010747664,0.0047670417,0.0007332688,0.00731906,0.0033950794,0.001438758,0.004323315],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046552686,0.00025296962,0.0113457525,0.0010654476,0.00037789432,0.0007975646,0.0009306276,0.04907229,0.031253926,0.044360455,0.06132199,0.7987556],"study_design_scores_gemma":[0.00007972266,0.00010893792,0.007492595,0.00024079993,0.00041086916,0.0014619634,0.0011827215,0.6358805,0.08782814,0.07398834,0.19107547,0.00024998572],"about_ca_topic_score_codex":0.003608834,"about_ca_topic_score_gemma":0.0069368742,"teacher_disagreement_score":0.0050259405,"about_ca_system_score_codex":0.00063277676,"about_ca_system_score_gemma":0.0016404409,"threshold_uncertainty_score":0.012945056},"labels":[],"label_agreement":null},{"id":"W3210704534","doi":"10.20944/preprints202110.0382.v1","title":"Employing Statistical Machine Reading for Inferring Key Concepts of a Research Field From a Body of Abstracts and Blog Posts","year":2021,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reading (process); Field (mathematics); Data science; Function (biology); Computer science; Comprehension; Conversation; Key (lock); Management science; Political science; Psychology; Engineering","score_opus":0.2271925651014114,"score_gpt":0.45563212320088736,"score_spread":0.22843955809947597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210704534","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28280184,0.012909271,0.6315567,0.012165881,0.0015839454,0.0021220052,0.019353509,0.006097578,0.03140928],"genre_scores_gemma":[0.47880304,0.0040183817,0.49687403,0.0011961479,0.0013046846,0.0016058863,0.012178934,0.00038790208,0.0036309734],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9860621,0.0072237533,0.0017214757,0.0023515953,0.0022958966,0.00034507553],"domain_scores_gemma":[0.7315613,0.22504035,0.017488465,0.010409795,0.014470303,0.0010297484],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.019748442,0.0016146055,0.0010742091,0.036317673,0.0016772529,0.006831852,0.0015836006,0.0017403365,0.004063988],"category_scores_gemma":[0.14457358,0.00054409343,0.0014025898,0.017956378,0.0022259112,0.009851145,0.002732766,0.0023402236,0.0037510507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007668639,0.00042395722,0.054838963,0.008539628,0.0005823614,0.0017381997,0.026078222,0.004285651,0.02213232,0.027559841,0.020339917,0.832714],"study_design_scores_gemma":[0.00022165696,0.001326802,0.114223376,0.0054390808,0.001168668,0.002663586,0.0533655,0.1959416,0.04289888,0.4040281,0.17773828,0.0009845189],"about_ca_topic_score_codex":0.0019001188,"about_ca_topic_score_gemma":0.004607971,"teacher_disagreement_score":0.96368235,"about_ca_system_score_codex":0.0015667172,"about_ca_system_score_gemma":0.0034245895,"threshold_uncertainty_score":0.10444093},"labels":[],"label_agreement":null},{"id":"W3211120245","doi":"10.5281/zenodo.5046661","title":"ExploreASL/ExploreASL: ExploreASL v1.7.0","year":2021,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science","score_opus":0.07513490076175912,"score_gpt":0.2492624769311891,"score_spread":0.17412757616942998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211120245","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009864433,0.0004604128,0.038506288,0.000574564,0.00040037784,0.00023981549,0.1681242,0.7798601,0.01084783],"genre_scores_gemma":[0.0076192897,0.0004241548,0.068713374,0.0011728394,0.00021325717,0.0022137307,0.314447,0.5902751,0.01492123],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99762005,0.00039959402,0.0002582548,0.0006481874,0.0007948047,0.00027912072],"domain_scores_gemma":[0.9962328,0.0012938434,0.00024736437,0.0008735921,0.0010242589,0.00032811356],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.004417896,0.0056345062,0.0026082564,0.004330926,0.0017290069,0.006616477,0.0067722425,0.002552229,0.31058824],"category_scores_gemma":[0.013391413,0.0033770283,0.004604759,0.003267679,0.0011238267,0.0059682503,0.0071822996,0.003920175,0.27139917],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035678752,0.000029895002,0.0006072501,0.0012697581,0.00012255978,0.00010205205,0.00019172407,0.0005934197,0.0029806553,0.0026167203,0.9698701,0.021259164],"study_design_scores_gemma":[0.00037388655,0.000085502834,0.0017736641,0.0003702552,0.00011118529,0.000256917,0.00014756541,0.0055493675,0.014010842,0.017336806,0.9597451,0.00023884598],"about_ca_topic_score_codex":0.004940266,"about_ca_topic_score_gemma":0.0071141734,"teacher_disagreement_score":0.31058824,"about_ca_system_score_codex":0.0017703132,"about_ca_system_score_gemma":0.0031683561,"threshold_uncertainty_score":0.9833622},"labels":[],"label_agreement":null},{"id":"W3211375313","doi":"10.18653/v1/2021.findings-emnlp.351","title":"Textual Time Travel: A Temporally Informed Approach to Theory of Mind","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Theory of mind; Mental model; Artificial intelligence; Prior probability; Question answering; False belief; Cognitive science; Psychology; Cognition","score_opus":0.030085367097814227,"score_gpt":0.24192704963633876,"score_spread":0.21184168253852453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211375313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059495796,0.0010305723,0.932265,0.0014620032,0.000099910634,0.00007330939,0.0010684128,0.0016497014,0.0028552928],"genre_scores_gemma":[0.80107605,0.0006289603,0.19399582,0.00024518414,0.00012793443,0.0001763575,0.0011432065,0.00012419002,0.0024822732],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9995771,0.00018054935,0.000026038371,0.00014361754,0.00004955769,0.000023108592],"domain_scores_gemma":[0.99731904,0.0016541128,0.00032837645,0.0004147146,0.00019607616,0.00008770273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011843459,0.00072092423,0.0004756806,0.0010395552,0.0004053952,0.0016023116,0.002356172,0.0013135062,0.0028752922],"category_scores_gemma":[0.007993695,0.00051323464,0.001267644,0.0011741286,0.0007580799,0.0047605103,0.0012106051,0.001899464,0.00047166125],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005813667,0.0002328355,0.010957641,0.0004806031,0.00034287816,0.00054526504,0.0021525398,0.65070486,0.007727745,0.108959734,0.0066082394,0.2107063],"study_design_scores_gemma":[0.000011130129,0.000035366487,0.00067500054,0.000015514168,0.00002923257,0.000060199964,0.000057966827,0.9452495,0.0007528229,0.05123444,0.0018609178,0.000017832432],"about_ca_topic_score_codex":0.011908556,"about_ca_topic_score_gemma":0.013748223,"teacher_disagreement_score":0.011908556,"about_ca_system_score_codex":0.0013665414,"about_ca_system_score_gemma":0.0008058705,"threshold_uncertainty_score":0.023678482},"labels":[],"label_agreement":null},{"id":"W3211439143","doi":"10.2196/27210","title":"A Question-and-Answer System to Extract Data From Free-Text Oncological Pathology Reports (CancerBERT Network): Development Study","year":2021,"lang":"en","type":"article","venue":"Journal of Medical Internet Research","topic":"Topic Modeling","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Alberta Health Services","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Terminology; SNOMED CT; Named-entity recognition; Pathology; Deep learning; Information retrieval; Medicine; Linguistics","score_opus":0.16471099525244154,"score_gpt":0.44931670532246165,"score_spread":0.2846057100700201,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211439143","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33760387,0.0033690792,0.45752522,0.002542256,0.0007489149,0.0039979266,0.01993209,0.15500854,0.01927199],"genre_scores_gemma":[0.33592883,0.0015956418,0.56457967,0.0014983586,0.00013134831,0.0019551506,0.062366627,0.0018676758,0.030076701],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989986,0.0002827119,0.00006683627,0.0002946874,0.00027951016,0.000077659075],"domain_scores_gemma":[0.9965288,0.0018998724,0.00013159298,0.0002995452,0.00093515037,0.00020497304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024903717,0.0010287659,0.00043719518,0.0010987552,0.00031142682,0.00083026814,0.0019628895,0.0012439336,0.007873648],"category_scores_gemma":[0.0054641464,0.00047656018,0.00066378526,0.0005662933,0.00042490204,0.00355539,0.0012807662,0.0014085122,0.0056082276],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001191094,0.0021538998,0.0096131,0.0017244158,0.00023807083,0.0014014503,0.0013720166,0.021025168,0.05611581,0.006394966,0.090058975,0.80871105],"study_design_scores_gemma":[0.000502148,0.0019104055,0.012357264,0.00028453657,0.0002193944,0.0020420633,0.00087277347,0.7507362,0.10817187,0.003526517,0.11917043,0.00020642317],"about_ca_topic_score_codex":0.010687988,"about_ca_topic_score_gemma":0.0099013215,"teacher_disagreement_score":0.010687988,"about_ca_system_score_codex":0.001127791,"about_ca_system_score_gemma":0.0017920611,"threshold_uncertainty_score":0.026340008},"labels":[],"label_agreement":null},{"id":"W3211439810","doi":"10.18653/v1/2021.emnlp-main.835","title":"Types of Out-of-Distribution Texts and How to Detect Them","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University","keywords":"Computer science; Categorization; Artificial intelligence; Calibration; Natural language processing; Data mining; Statistics; Mathematics","score_opus":0.0736749441821645,"score_gpt":0.3899331424867225,"score_spread":0.31625819830455804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211439810","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6693535,0.008470272,0.26912695,0.0072822515,0.0010589169,0.0008162445,0.015986644,0.007240902,0.02066438],"genre_scores_gemma":[0.82513046,0.0016833829,0.15195207,0.000688033,0.00039259865,0.00038223053,0.014468663,0.0009002696,0.004402319],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9901248,0.003549391,0.0011089261,0.0022155228,0.0026423715,0.000358915],"domain_scores_gemma":[0.938959,0.04427333,0.005402951,0.0051927674,0.0051868227,0.0009851063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057848627,0.0008816297,0.0009364073,0.006984505,0.0010901827,0.0044010617,0.0016023337,0.0022039886,0.0028611177],"category_scores_gemma":[0.07817785,0.00055587,0.0007307596,0.003228913,0.0013248514,0.0074931486,0.002227371,0.0019687517,0.0027951356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010737324,0.00033816983,0.30419278,0.004581575,0.00041786605,0.0025023734,0.008854931,0.007849592,0.025042387,0.023959707,0.06861564,0.5525713],"study_design_scores_gemma":[0.00021220402,0.00035816757,0.24583809,0.0024623813,0.00046283758,0.016059548,0.01956075,0.3659904,0.071207255,0.09008551,0.18728623,0.00047668652],"about_ca_topic_score_codex":0.0014084377,"about_ca_topic_score_gemma":0.0022756108,"teacher_disagreement_score":0.006984505,"about_ca_system_score_codex":0.00075715006,"about_ca_system_score_gemma":0.00072160165,"threshold_uncertainty_score":0.030593693},"labels":[],"label_agreement":null},{"id":"W3211673252","doi":"10.5281/zenodo.3982680","title":"Related codes for \"Why Do Masked Neural Language Models Still Need Commonsense Knowledge to Handle Semantic Variations in Question Answering?\"","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Question answering; Commonsense knowledge; Computer science; Natural language processing; Artificial intelligence; Commonsense reasoning; Knowledge-based systems","score_opus":0.04508286160316482,"score_gpt":0.26637454174296293,"score_spread":0.2212916801397981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211673252","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008951596,0.0044571976,0.35737708,0.29028755,0.1841348,0.00043942634,0.0114132855,0.005390952,0.13754815],"genre_scores_gemma":[0.3709563,0.0049493345,0.12910984,0.089543924,0.077151895,0.0012883085,0.019549116,0.008224454,0.29922682],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99635625,0.0008376747,0.0002922831,0.00082281965,0.0013630728,0.00032791865],"domain_scores_gemma":[0.97820866,0.0077118766,0.0010271505,0.003550599,0.008900271,0.0006014915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031277537,0.0010378141,0.00086578337,0.0016026283,0.0020506715,0.0035590117,0.0022516607,0.004643602,0.09455059],"category_scores_gemma":[0.03920053,0.00051256805,0.0015748108,0.0015003426,0.0025164988,0.0066674724,0.0029045828,0.0036786275,0.03150907],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018064649,0.000073301126,0.0005047124,0.0002863865,0.000036968693,0.0001619192,0.0002663263,0.001053859,0.002250972,0.3355287,0.6248922,0.034764025],"study_design_scores_gemma":[0.00010887189,0.000059010792,0.0019228067,0.00017023925,0.0000417385,0.00047660252,0.0002632417,0.031194577,0.009370142,0.5482721,0.40797067,0.00014998093],"about_ca_topic_score_codex":0.005278619,"about_ca_topic_score_gemma":0.0025839047,"teacher_disagreement_score":0.09455059,"about_ca_system_score_codex":0.003271628,"about_ca_system_score_gemma":0.0022047027,"threshold_uncertainty_score":0.3163032},"labels":[],"label_agreement":null},{"id":"W3211782300","doi":"10.48550/arxiv.2012.15495","title":"Towards Zero-Shot Knowledge Distillation for Natural Language Processing","year":2020,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Distillation; Task (project management); Benchmark (surveying); Artificial intelligence; Knowledge transfer; Natural language processing; Variety (cybernetics); Domain knowledge; Machine learning; Transfer of learning; Domain (mathematical analysis); Shot (pellet); Knowledge management; Mathematics; Engineering","score_opus":0.09527768866442518,"score_gpt":0.21964052195364003,"score_spread":0.12436283328921485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211782300","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052122306,0.0013329961,0.932523,0.0015960814,0.000142315,0.00011351253,0.0005897223,0.006052371,0.005527645],"genre_scores_gemma":[0.5955835,0.0007924551,0.3899967,0.0009699405,0.0002159253,0.0002709607,0.0027682781,0.0008521605,0.0085499985],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985514,0.0005511214,0.000057568486,0.00033833494,0.00036752905,0.00013394447],"domain_scores_gemma":[0.99673057,0.0020894478,0.00012307696,0.0006708782,0.0002521943,0.00013384478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024421888,0.0013418857,0.0011285293,0.0009144751,0.00084272266,0.001686172,0.0023225625,0.0020480568,0.0032537314],"category_scores_gemma":[0.008751475,0.00053925446,0.0009627036,0.0009587877,0.0020471758,0.0051594027,0.0047490476,0.004868771,0.0017337814],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075813464,0.0004982119,0.0014858249,0.0006650536,0.00015975478,0.000307487,0.0006650162,0.45079142,0.013542028,0.09692349,0.021610381,0.41259328],"study_design_scores_gemma":[0.000028638779,0.00006382809,0.00008908231,0.000023003018,0.0000102676895,0.00004623795,0.000046171328,0.9294652,0.0048159095,0.06304222,0.0023581502,0.000011290733],"about_ca_topic_score_codex":0.003486001,"about_ca_topic_score_gemma":0.0056907022,"teacher_disagreement_score":0.003486001,"about_ca_system_score_codex":0.0011511019,"about_ca_system_score_gemma":0.0017501635,"threshold_uncertainty_score":0.0129157305},"labels":[],"label_agreement":null},{"id":"W3212619741","doi":"10.18653/v1/2021.eval4nlp-1.6","title":"Trainable Ranking Models to Evaluate the Semantic Accuracy of Data-to-Text Neural Generator","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Metric (unit); Ranking (information retrieval); Generalization; Generator (circuit theory); Table (database); Artificial intelligence; Inference; Embedding; Natural language processing; Artificial neural network; Machine learning; Data mining; Mathematics","score_opus":0.18353008287572417,"score_gpt":0.339522803652108,"score_spread":0.15599272077638385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3212619741","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20729694,0.004116499,0.7452681,0.0016311903,0.0008386048,0.0007510301,0.0072735506,0.021065336,0.01175878],"genre_scores_gemma":[0.79755586,0.00058043224,0.17503768,0.000574566,0.0002290266,0.0005716578,0.019987637,0.0012210531,0.004241984],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99283653,0.0034085233,0.00057312974,0.0013298199,0.0015168362,0.00033522802],"domain_scores_gemma":[0.97317356,0.01826717,0.001245448,0.004275791,0.0025312696,0.0005067193],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012186468,0.0027487578,0.001185062,0.003451145,0.00073750905,0.0022701912,0.0025949702,0.0027878587,0.0034094884],"category_scores_gemma":[0.052948218,0.00038274517,0.0011145582,0.002184704,0.0014550976,0.0046687173,0.0024601512,0.002827056,0.0022214402],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010589084,0.0006330718,0.021609632,0.0012894702,0.0008432943,0.0004265743,0.0004240324,0.37922627,0.013493508,0.01005295,0.025823092,0.5451192],"study_design_scores_gemma":[0.00006374966,0.0004171617,0.0033399952,0.000053935073,0.00009512615,0.0002186092,0.00015655675,0.9651441,0.01544782,0.0116457455,0.0033645304,0.000052578456],"about_ca_topic_score_codex":0.0036589683,"about_ca_topic_score_gemma":0.0052651186,"teacher_disagreement_score":0.012186468,"about_ca_system_score_codex":0.0015907836,"about_ca_system_score_gemma":0.0013992551,"threshold_uncertainty_score":0.06444895},"labels":[],"label_agreement":null},{"id":"W3212682848","doi":"10.48550/arxiv.2112.00578","title":"Systematic Generalization with Edge Transformers","year":2021,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Transformer; Computer science; Parsing; Unification; Artificial intelligence; Theoretical computer science; Programming language; Engineering; Electrical engineering","score_opus":0.04358110865544729,"score_gpt":0.15750120961687417,"score_spread":0.11392010096142688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3212682848","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039546043,0.0005042467,0.94562817,0.00077752565,0.00007137552,0.00008286183,0.0005378549,0.0047208564,0.00813108],"genre_scores_gemma":[0.7642414,0.00071761204,0.22582158,0.0006682271,0.00008684419,0.00017183559,0.0016434342,0.00076712744,0.0058819046],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990257,0.00024600682,0.00006345628,0.0003875571,0.00019418083,0.00008310664],"domain_scores_gemma":[0.9963696,0.0016550268,0.00017154784,0.0014287006,0.00027434947,0.0001007897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015800241,0.0006812944,0.00077770714,0.0009673619,0.0005025271,0.0015242405,0.0021069702,0.0010054316,0.00591977],"category_scores_gemma":[0.008129583,0.00043747164,0.0013465239,0.0011947955,0.0015169336,0.009155904,0.0029648931,0.0024575326,0.0011889826],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032060174,0.00015529338,0.002982206,0.000318388,0.0001275209,0.00025007475,0.0005636535,0.20744574,0.0077975057,0.38544673,0.011780596,0.38281167],"study_design_scores_gemma":[0.000017245977,0.000034435536,0.00021101427,0.000018799246,0.00003895495,0.00008716658,0.000045423927,0.61613744,0.002497311,0.37651786,0.004382179,0.000012126008],"about_ca_topic_score_codex":0.0036254402,"about_ca_topic_score_gemma":0.005734633,"teacher_disagreement_score":0.00591977,"about_ca_system_score_codex":0.0009955732,"about_ca_system_score_gemma":0.0011520847,"threshold_uncertainty_score":0.019803643},"labels":[],"label_agreement":null},{"id":"W3212805212","doi":"10.5281/zenodo.4244655","title":"2020 Programming Historian Deposit release","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; University of Waterloo","funders":"","keywords":"Geology; Computer science","score_opus":0.03731586346728265,"score_gpt":0.2255656860574279,"score_spread":0.18824982259014528,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3212805212","genre_codex":"other","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018462747,0.0010411838,0.018855406,0.005943541,0.0036353853,0.00031052495,0.22946234,0.10349949,0.63540596],"genre_scores_gemma":[0.0058153733,0.0009790979,0.012090372,0.00096130476,0.00066579314,0.00035349777,0.16523357,0.06351431,0.75038666],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.997366,0.00027161473,0.00017042141,0.0003175537,0.0016950829,0.00017931231],"domain_scores_gemma":[0.98978513,0.0010111416,0.0003732039,0.0024920986,0.0054132547,0.0009252415],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0043940153,0.0009104416,0.0009038489,0.005906325,0.0017597801,0.008138273,0.0019592515,0.0012961328,0.5800481],"category_scores_gemma":[0.018527089,0.0013031401,0.000728361,0.0070673483,0.0005634081,0.0070725763,0.0046824645,0.0019696963,0.60286105],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048051515,0.00001051595,0.000119354285,0.00010823713,0.000002788047,0.000018993607,0.00007804338,0.00004270136,0.00013627071,0.003709996,0.96755844,0.028166598],"study_design_scores_gemma":[0.000006935036,0.0000028188233,0.0001346999,0.000041332813,0.0000017012582,0.000024909514,0.000024121888,0.00005880603,0.00017202832,0.000990501,0.9985354,0.000006644599],"about_ca_topic_score_codex":0.007940676,"about_ca_topic_score_gemma":0.008759169,"teacher_disagreement_score":0.5800481,"about_ca_system_score_codex":0.0027069887,"about_ca_system_score_gemma":0.0047428873,"threshold_uncertainty_score":0.59901047},"labels":[],"label_agreement":null},{"id":"W3213180921","doi":"10.18653/v1/2021.emnlp-main.603","title":"Universal-KD: Attention-based Output-Grounded Intermediate Layer Knowledge Distillation","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Layer (electronics); Distillation; Matching (statistics); Interpretability; Projection (relational algebra); Base (topology); Space (punctuation); Architecture; Artificial intelligence; Deep learning; Computer architecture; Machine learning; Algorithm; Mathematics; Nanotechnology; Chromatography; Operating system; Materials science; Chemistry","score_opus":0.07139261601228206,"score_gpt":0.3974343703737435,"score_spread":0.32604175436146143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3213180921","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017353443,0.0005491757,0.9684529,0.00028498739,0.00015982534,0.00010035972,0.00037450745,0.009335288,0.0033895078],"genre_scores_gemma":[0.5222674,0.00040823076,0.46204996,0.00061678817,0.000105534804,0.00026222182,0.0024816853,0.0007841099,0.011024045],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99910957,0.00014027383,0.00006651468,0.00033748572,0.00017817947,0.0001680753],"domain_scores_gemma":[0.9987871,0.00041302346,0.00008296138,0.00041826963,0.0002216544,0.00007696138],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011065565,0.0014869481,0.0012436754,0.0012074496,0.0007569246,0.0015306475,0.0035293717,0.0014901153,0.0072759916],"category_scores_gemma":[0.0040924614,0.00066296116,0.0013472692,0.0014672427,0.0010898384,0.0053373547,0.0043886886,0.0030357013,0.002226051],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033267765,0.00023474319,0.0010606756,0.00038991056,0.00013649529,0.00022077038,0.0004677069,0.11240563,0.015390086,0.026427519,0.0129435165,0.8299904],"study_design_scores_gemma":[0.000040734674,0.00008623925,0.00032448946,0.000038821803,0.00006153434,0.00009096831,0.00008379138,0.9381472,0.017062612,0.03694777,0.0070834947,0.000032338427],"about_ca_topic_score_codex":0.007131571,"about_ca_topic_score_gemma":0.010524195,"teacher_disagreement_score":0.0072759916,"about_ca_system_score_codex":0.0011972333,"about_ca_system_score_gemma":0.0020891977,"threshold_uncertainty_score":0.02434063},"labels":[],"label_agreement":null},{"id":"W3213418658","doi":"10.18653/v1/2021.mrl-1.11","title":"Small Data? No Problem! Exploring the Viability of Pretrained Multilingual Language Models for Low-resourced Languages","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":114,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Language model; Second-generation programming language; Code (set theory); Resource (disambiguation); Variety (cybernetics); Training set; Programming language; Set (abstract data type); Fifth-generation programming language; Programming paradigm","score_opus":0.12160079997991133,"score_gpt":0.29630273869832713,"score_spread":0.1747019387184158,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3213418658","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18998447,0.0077613685,0.6902302,0.040193804,0.0027750763,0.00047778958,0.015577209,0.028819619,0.024180476],"genre_scores_gemma":[0.5542475,0.0028098717,0.38023004,0.0060593723,0.0008673729,0.0008676629,0.038359504,0.0052268067,0.011331939],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973443,0.0013014117,0.00014758554,0.0007079163,0.00032753462,0.00017133975],"domain_scores_gemma":[0.9907354,0.0052097565,0.00019958394,0.0025276172,0.00084578135,0.00048187608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062351134,0.0019193862,0.0013359281,0.0010026962,0.0012618648,0.0033624235,0.0035217497,0.0017239306,0.007594994],"category_scores_gemma":[0.023897836,0.0011374777,0.0018993121,0.0014869523,0.0018984835,0.014737093,0.0044770045,0.006245129,0.006235125],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021324276,0.00066635816,0.024504026,0.0015750908,0.0012546562,0.00097626523,0.0021012018,0.17899278,0.015678572,0.05008725,0.24107696,0.48095438],"study_design_scores_gemma":[0.0004127255,0.0003914182,0.0035618497,0.00041987118,0.00027302184,0.00046982546,0.0016506455,0.7446854,0.013718691,0.14921407,0.08505466,0.00014773017],"about_ca_topic_score_codex":0.011244131,"about_ca_topic_score_gemma":0.017907035,"teacher_disagreement_score":0.011244131,"about_ca_system_score_codex":0.0011195927,"about_ca_system_score_gemma":0.0016071092,"threshold_uncertainty_score":0.03297478},"labels":[],"label_agreement":null},{"id":"W3213458975","doi":"10.1162/tacl_a_00419","title":"<scp>ParsiNLU</scp>: A Suite of Language Understanding Challenges for Persian","year":2021,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); International Medias Data Services (Canada)","funders":"","keywords":"Computer science; Natural language understanding; Natural language processing; Suite; Benchmark (surveying); Artificial intelligence; Persian; Natural language; Linguistics","score_opus":0.053844189722149584,"score_gpt":0.28161477226806203,"score_spread":0.22777058254591245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3213458975","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6244645,0.010143142,0.048771054,0.016925117,0.0021807386,0.0005722797,0.18911153,0.048113808,0.059717897],"genre_scores_gemma":[0.60130185,0.0011557904,0.043444943,0.0017260293,0.00036799506,0.0004199042,0.34032878,0.0017359968,0.009518782],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99752575,0.0010414777,0.00019295223,0.00062146573,0.00043827758,0.0001799329],"domain_scores_gemma":[0.9946214,0.0025354326,0.00020100563,0.0010353022,0.0012672788,0.00033962476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020708523,0.0019299705,0.0008625564,0.0029155158,0.0022568882,0.0021855503,0.0016472975,0.0016759759,0.0078751715],"category_scores_gemma":[0.009059928,0.0002945137,0.0006941662,0.0029611925,0.0012440592,0.0039162855,0.0028468135,0.0018288923,0.005960732],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073087745,0.00044177094,0.014574258,0.0017320957,0.00019055043,0.003280389,0.0035336134,0.019816294,0.00939104,0.00863805,0.61807954,0.3195915],"study_design_scores_gemma":[0.00041758188,0.000474011,0.05564914,0.0007659082,0.0001663215,0.0059808455,0.017981257,0.1691344,0.057866875,0.05207434,0.63911057,0.0003787563],"about_ca_topic_score_codex":0.025297292,"about_ca_topic_score_gemma":0.032802656,"teacher_disagreement_score":0.025297292,"about_ca_system_score_codex":0.0016360655,"about_ca_system_score_gemma":0.0024148687,"threshold_uncertainty_score":0.05030006},"labels":[],"label_agreement":null},{"id":"W3213719910","doi":"10.18653/v1/2021.sustainlp-1.8","title":"Learning to Rank in the Age of Muppets: Effectiveness–Efficiency Tradeoffs in Multi-Stage Ranking","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; University of Waterloo; Compute Canada","keywords":"Transformer; Computer science; Inference; Machine learning; Artificial intelligence; Ranking (information retrieval); Learning to rank; Recall; Engineering; Voltage","score_opus":0.049216727700560924,"score_gpt":0.30256823670876254,"score_spread":0.25335150900820164,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3213719910","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14861654,0.008619179,0.8198367,0.0034236223,0.00032773733,0.00025750464,0.0006238531,0.0053282343,0.012966732],"genre_scores_gemma":[0.6845441,0.0019608485,0.3008749,0.00041670448,0.0005594051,0.00010866789,0.0008815965,0.00056601455,0.010087784],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956201,0.002087778,0.00027717973,0.00055695925,0.0010473011,0.0004107183],"domain_scores_gemma":[0.97185326,0.019532006,0.0009453283,0.0048801852,0.002073337,0.0007158205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0109683415,0.0012153527,0.0021802492,0.0022736837,0.0011982779,0.0031939705,0.0022930477,0.0021484187,0.004113747],"category_scores_gemma":[0.03747265,0.0007792409,0.0007805121,0.0019104667,0.001183242,0.009069952,0.0018858245,0.0024203171,0.0037776956],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014708419,0.0005605852,0.009157177,0.00053518097,0.000231907,0.0003970471,0.00051923364,0.17606185,0.012223254,0.056435063,0.023114422,0.7192935],"study_design_scores_gemma":[0.00009799196,0.0007013599,0.0015907402,0.00004450445,0.00009178081,0.000653176,0.0001939845,0.9294857,0.010532954,0.05136373,0.005184548,0.000059620757],"about_ca_topic_score_codex":0.0030027353,"about_ca_topic_score_gemma":0.0060331756,"teacher_disagreement_score":0.0109683415,"about_ca_system_score_codex":0.0008407583,"about_ca_system_score_gemma":0.0011669025,"threshold_uncertainty_score":0.058006823},"labels":[],"label_agreement":null},{"id":"W3213905624","doi":"10.18653/v1/2021.blackboxnlp-1.35","title":"Interacting Knowledge Sources, Inspection and Analysis: Case-studies on Biomedical text processing","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Encapsulation (networking); Transparency (behavior)","score_opus":0.04572935156477141,"score_gpt":0.3320811119229192,"score_spread":0.2863517603581478,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3213905624","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5307596,0.0034146772,0.45152995,0.0030330739,0.000051797815,0.0009125496,0.0009843189,0.0016312435,0.0076827426],"genre_scores_gemma":[0.81822413,0.0011621602,0.17686689,0.00030718782,0.00005791794,0.00024303106,0.00096903916,0.00040800776,0.0017616162],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98560196,0.009454016,0.0008094399,0.0012629629,0.0025271801,0.00034454613],"domain_scores_gemma":[0.84292513,0.1418159,0.0043190904,0.007410821,0.0030056434,0.00052332063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014605317,0.00090654625,0.0010835897,0.004678213,0.0020741522,0.003935568,0.0029580772,0.0048150048,0.002457334],"category_scores_gemma":[0.07756454,0.0007010585,0.0014196545,0.0052381186,0.003034048,0.006074735,0.0035646816,0.0020697294,0.0007169579],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030979277,0.003247481,0.10229838,0.0070457426,0.0009978607,0.026393687,0.058396623,0.059417766,0.026739512,0.04297331,0.010840688,0.65855116],"study_design_scores_gemma":[0.0009960999,0.0018940722,0.06874167,0.0019788893,0.0023742449,0.032381732,0.028866334,0.45498857,0.151945,0.16777563,0.08747999,0.00057779904],"about_ca_topic_score_codex":0.0046240436,"about_ca_topic_score_gemma":0.0046497597,"teacher_disagreement_score":0.014605317,"about_ca_system_score_codex":0.0011403549,"about_ca_system_score_gemma":0.0013732401,"threshold_uncertainty_score":0.07724118},"labels":[],"label_agreement":null},{"id":"W3214002337","doi":"","title":"Combiner: Full Attention Transformer with Sparse Computation Cost","year":2021,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Computation; Quadratic equation; Leverage (statistics); Transformer; Factorization; Autoregressive model; Theoretical computer science; Algorithm; Artificial intelligence; Mathematics; Voltage","score_opus":0.06034433982791471,"score_gpt":0.17640755943918549,"score_spread":0.11606321961127078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214002337","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0091995895,0.0003464901,0.97718245,0.00018154392,0.0000740178,0.000094882525,0.0002540983,0.009788397,0.0028785602],"genre_scores_gemma":[0.4503664,0.00071062415,0.5295938,0.0006889862,0.00018997122,0.00039642933,0.001580477,0.0013352437,0.015138059],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994584,0.00008741072,0.00002491961,0.00017940681,0.00016951394,0.000080366335],"domain_scores_gemma":[0.9994673,0.0002135096,0.000038888735,0.00014499052,0.000086598855,0.000048778005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007423826,0.0015453427,0.0009687781,0.00071616616,0.00035987506,0.0010907192,0.0022528223,0.0008639214,0.010345948],"category_scores_gemma":[0.002607082,0.0006316253,0.001035841,0.0007420085,0.0006598488,0.0031382665,0.0022527142,0.0016846291,0.0038753634],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006081452,0.00020479005,0.0010718893,0.0003254866,0.00013088711,0.00031207313,0.00023862188,0.08173067,0.04097829,0.042320084,0.02423278,0.80784637],"study_design_scores_gemma":[0.000066935674,0.00013451683,0.00027379137,0.000016913633,0.00005387895,0.00018026984,0.000040690564,0.9386196,0.015502724,0.038766984,0.006320031,0.00002356674],"about_ca_topic_score_codex":0.0073688175,"about_ca_topic_score_gemma":0.013261821,"teacher_disagreement_score":0.010345948,"about_ca_system_score_codex":0.001026302,"about_ca_system_score_gemma":0.0014388305,"threshold_uncertainty_score":0.03461063},"labels":[],"label_agreement":null},{"id":"W3214038065","doi":"10.18653/v1/2021.findings-emnlp.26","title":"Multi-Task Dense Retrieval via Model Uncertainty Fusion for Open-Domain Question Answering","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Computer science; Benchmark (surveying); Task (project management); Question answering; Information retrieval; Joint (building); Set (abstract data type); Domain (mathematical analysis); Aggregate (composite); Artificial intelligence; Training set; Code (set theory); Language model; Natural language processing","score_opus":0.041967072789004724,"score_gpt":0.3043780166640182,"score_spread":0.26241094387501346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214038065","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02643147,0.002264614,0.9629656,0.000487031,0.00007973677,0.00015168004,0.00063544814,0.0058209733,0.0011634344],"genre_scores_gemma":[0.6399652,0.00095714023,0.34626332,0.0009812133,0.0004641749,0.00038377658,0.0065340744,0.0005293789,0.003921809],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997095,0.00097246666,0.00018564053,0.00094837305,0.0005521159,0.0002463871],"domain_scores_gemma":[0.99455935,0.0033122632,0.00029988863,0.0010052224,0.00066008576,0.00016323286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00431908,0.0023814258,0.002648175,0.0027730286,0.00076085585,0.0016310635,0.0027675563,0.0024136282,0.0023316282],"category_scores_gemma":[0.011022555,0.0009398226,0.0023149953,0.0024592236,0.0010344298,0.0043729427,0.0033883902,0.0028790243,0.0015384868],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065847934,0.0005851006,0.0033000116,0.0006629489,0.0006649936,0.0004154936,0.0007615098,0.35550538,0.0153919,0.00728374,0.017282104,0.59748834],"study_design_scores_gemma":[0.000035574187,0.000112783484,0.0005053284,0.00001862146,0.00007255757,0.00010086876,0.00006836301,0.9832697,0.0029229869,0.0110879075,0.001771797,0.000033587203],"about_ca_topic_score_codex":0.009815631,"about_ca_topic_score_gemma":0.009218999,"teacher_disagreement_score":0.009815631,"about_ca_system_score_codex":0.0012619669,"about_ca_system_score_gemma":0.001383628,"threshold_uncertainty_score":0.022841811},"labels":[],"label_agreement":null},{"id":"W3214151748","doi":"10.18653/v1/2021.mrl-1.12","title":"Mr. TyDi: A Multi-lingual Benchmark for Dense Retrieval","year":2021,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Computer science; Benchmark (surveying); Ranking (information retrieval); Relevance (law); Artificial intelligence; Natural language processing; Information retrieval; Representation (politics); Point (geometry); Resource (disambiguation); Machine learning; Mathematics; Geography","score_opus":0.050235439456535906,"score_gpt":0.30067412928951526,"score_spread":0.25043868983297934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214151748","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21588124,0.02418074,0.061343603,0.0048731165,0.004103779,0.0036514266,0.57289255,0.045014836,0.068058684],"genre_scores_gemma":[0.11969156,0.0016688245,0.06696429,0.0009908493,0.0004750916,0.0011537687,0.7935824,0.0014053494,0.014067795],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9929646,0.002082522,0.00091706373,0.0012785562,0.0020633365,0.000693954],"domain_scores_gemma":[0.9911901,0.0023667112,0.00051810715,0.0028718652,0.0023742232,0.00067907263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005776885,0.0029680098,0.0020897714,0.007847357,0.0027238568,0.00355911,0.0041434844,0.0029853159,0.012203165],"category_scores_gemma":[0.019203907,0.00058874063,0.002122662,0.008265987,0.0015101222,0.005421742,0.0046685333,0.0023878329,0.012880997],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001461351,0.001457493,0.0063065104,0.004622561,0.0007308591,0.00057151326,0.00053897424,0.013862304,0.011792494,0.005309647,0.7420002,0.21134605],"study_design_scores_gemma":[0.0016195881,0.002583578,0.032403067,0.00072651735,0.0006859622,0.0042994004,0.0028572793,0.21307515,0.037825026,0.014051851,0.68912375,0.00074886857],"about_ca_topic_score_codex":0.046968736,"about_ca_topic_score_gemma":0.07484346,"teacher_disagreement_score":0.046968736,"about_ca_system_score_codex":0.0025703632,"about_ca_system_score_gemma":0.0035213965,"threshold_uncertainty_score":0.0933907},"labels":[],"label_agreement":null},{"id":"W3214455632","doi":"10.18653/v1/2021.emnlp-main.77","title":"Contextualized Query Embeddings for Conversational Search","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Leverage (statistics); Inference; Security token; Relevance (law); Pipeline (software); Information retrieval; Query expansion; Query language; Artificial intelligence","score_opus":0.092333595524727,"score_gpt":0.43953950920379287,"score_spread":0.34720591367906584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214455632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027039591,0.0012971435,0.96456516,0.0005952728,0.00007895151,0.0001178203,0.0012272894,0.0025866535,0.0024922525],"genre_scores_gemma":[0.7218541,0.0009861687,0.26397416,0.00046500043,0.00019568762,0.00043441667,0.0040412676,0.00042863836,0.0076205716],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992386,0.00026611466,0.000051209132,0.00023581735,0.00012628762,0.00008193574],"domain_scores_gemma":[0.9986927,0.0006620416,0.000108426066,0.00028351645,0.00019791338,0.000055396045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009613278,0.0008798506,0.00079085224,0.00096645224,0.0004169227,0.00109613,0.0016131026,0.0013192791,0.0042360625],"category_scores_gemma":[0.0069629415,0.00048180792,0.0008611562,0.0011284112,0.00068074086,0.004279525,0.0014628662,0.001848169,0.0018623141],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065787585,0.0004651977,0.0035383173,0.00059863273,0.00018583833,0.00031479806,0.0010712699,0.43806502,0.018018061,0.12175751,0.024133425,0.391194],"study_design_scores_gemma":[0.000014445077,0.000053804393,0.00024479596,0.000013639749,0.000020995667,0.00006621602,0.000047213587,0.96408737,0.001342468,0.0311892,0.0029047893,0.000015136694],"about_ca_topic_score_codex":0.008016057,"about_ca_topic_score_gemma":0.010511986,"teacher_disagreement_score":0.008016057,"about_ca_system_score_codex":0.0011765318,"about_ca_system_score_gemma":0.0010077246,"threshold_uncertainty_score":0.015938759},"labels":[],"label_agreement":null},{"id":"W3215214493","doi":"10.1145/3473973","title":"Graph Neural Collaborative Topic Model for Citation Recommendation","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Citation; Pairwise comparison; Graph; Information retrieval; Knowledge graph; Artificial neural network; Information overload; Data science; Artificial intelligence; Machine learning; Theoretical computer science; World Wide Web","score_opus":0.042966983719350094,"score_gpt":0.276904698131039,"score_spread":0.23393771441168892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3215214493","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.065534115,0.0060940827,0.9143344,0.0017466584,0.0002816149,0.00012023316,0.0018982504,0.0021035986,0.0078869155],"genre_scores_gemma":[0.82381433,0.0053486656,0.14187306,0.0006147679,0.0007119492,0.0004542647,0.0044454373,0.0002382417,0.022499233],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993394,0.00020408246,0.0000399118,0.00021039775,0.00013684941,0.00006933927],"domain_scores_gemma":[0.99862576,0.0008141765,0.00013857035,0.00011153825,0.00025504324,0.000054964832],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0013106195,0.0008877128,0.0013577092,0.002563164,0.0005530996,0.0012364961,0.0024274047,0.0022070622,0.00351356],"category_scores_gemma":[0.0054871584,0.0004851269,0.001255994,0.00431988,0.0005256552,0.0030792612,0.0006452528,0.0017571637,0.0016934122],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029261853,0.00028369066,0.0040402487,0.00035685478,0.00032315144,0.00017680626,0.00026565048,0.7455272,0.0021528222,0.052448057,0.014211293,0.17992152],"study_design_scores_gemma":[0.000010754821,0.000012855397,0.00026582432,0.0000074230593,0.000027440752,0.000017333092,0.0000064626447,0.98723674,0.00013093237,0.011417237,0.0008592271,0.000007758222],"about_ca_topic_score_codex":0.026032526,"about_ca_topic_score_gemma":0.027777351,"teacher_disagreement_score":0.9974368,"about_ca_system_score_codex":0.0015189846,"about_ca_system_score_gemma":0.0010694513,"threshold_uncertainty_score":0.051762044},"labels":[],"label_agreement":null},{"id":"W3217119495","doi":"10.1109/access.2021.3129786","title":"A Survey of Automatic Text Summarization: Progress, Process and Challenges","year":2021,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":146,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Automatic summarization; Computer science; Workflow; Taxonomy (biology); Information retrieval; Domain (mathematical analysis); Text graph; Process (computing); The Internet; Feature extraction; Data science; Artificial intelligence; World Wide Web; Database","score_opus":0.08943312503184098,"score_gpt":0.32881871083961284,"score_spread":0.23938558580777186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3217119495","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0201966,0.35527852,0.5826624,0.0076139523,0.0017296589,0.00058942067,0.004307852,0.011838463,0.015783224],"genre_scores_gemma":[0.104032315,0.35802925,0.49432686,0.002489408,0.005653371,0.0007447583,0.01862855,0.0022380229,0.013857561],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973521,0.0007518843,0.00032320767,0.00058828533,0.00086588366,0.00011870602],"domain_scores_gemma":[0.99044985,0.005258782,0.00063443085,0.0008806717,0.0025993993,0.00017685929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031915335,0.0016355078,0.0016225707,0.006435175,0.0009030075,0.0035869924,0.0018014361,0.001133537,0.0040036174],"category_scores_gemma":[0.011464068,0.00061473553,0.0010349326,0.0072872248,0.0007933763,0.006973457,0.0013189587,0.0014953705,0.0057495544],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000092867915,0.00006457427,0.00090715906,0.0042456244,0.000068178066,0.00005752311,0.00042775765,0.00245451,0.006437891,0.00526428,0.030983787,0.94899577],"study_design_scores_gemma":[0.00006715802,0.0007240675,0.008548386,0.0040793056,0.0004901664,0.0011368922,0.0031064367,0.09542143,0.052232143,0.03746977,0.7964034,0.0003209349],"about_ca_topic_score_codex":0.0018231292,"about_ca_topic_score_gemma":0.0020006818,"teacher_disagreement_score":0.006435175,"about_ca_system_score_codex":0.0008163887,"about_ca_system_score_gemma":0.0016780285,"threshold_uncertainty_score":0.016878605},"labels":[],"label_agreement":null},{"id":"W3217492267","doi":"10.1007/978-3-030-99739-7_18","title":"Less is Less: When are Snippets Insufficient for Human vs Machine Relevance Estimation?","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Snippet; Computer science; Information retrieval; Relevance (law); Ranking (information retrieval); Search engine; Document retrieval; Query expansion; Language model; Range (aeronautics); Artificial intelligence; Natural language processing","score_opus":0.03607754583885837,"score_gpt":0.27603009777235954,"score_spread":0.23995255193350118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3217492267","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.711716,0.02762343,0.16470729,0.020850407,0.003379047,0.0006108157,0.0076444284,0.010892748,0.052575774],"genre_scores_gemma":[0.92917657,0.0027186568,0.04800677,0.0024639,0.0019665721,0.00015865959,0.005981333,0.002254799,0.007272759],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99301094,0.0025632514,0.0006555503,0.0012136868,0.0021745323,0.00038193958],"domain_scores_gemma":[0.9400612,0.044486895,0.0027107084,0.0038865302,0.0077074417,0.0011472204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008126549,0.000887038,0.0013831222,0.0036997474,0.0012741758,0.0037257862,0.0011933686,0.0027148318,0.0076085445],"category_scores_gemma":[0.09465191,0.0005829975,0.00069648813,0.0022247338,0.0014196599,0.014073257,0.0015270459,0.0023028618,0.006082228],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038934483,0.00054946356,0.08007272,0.0016573101,0.000391297,0.0007938748,0.0016074156,0.00277812,0.019923933,0.0072360346,0.08092405,0.8001724],"study_design_scores_gemma":[0.0005257249,0.0027852282,0.28032798,0.003185168,0.00258751,0.008564576,0.014012949,0.20057954,0.065328516,0.21441633,0.2069804,0.00070617156],"about_ca_topic_score_codex":0.0032094258,"about_ca_topic_score_gemma":0.005802188,"teacher_disagreement_score":0.008126549,"about_ca_system_score_codex":0.00068660383,"about_ca_system_score_gemma":0.0013065203,"threshold_uncertainty_score":0.04297775},"labels":[],"label_agreement":null},{"id":"W3217747958","doi":"10.18280/isi.260506","title":"Impact of Using Bidirectional Encoder Representations from Transformers (BERT) Models for Arabic Dialogue Acts Identification","year":2021,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Transformer; Natural language processing; Encoder; Artificial intelligence; Representation (politics); Utterance; Language model; Arabic; Identification (biology); Task (project management); Speech recognition; Linguistics","score_opus":0.0470122826845457,"score_gpt":0.2962686801892959,"score_spread":0.2492563975047502,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3217747958","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55199397,0.0067007067,0.40655568,0.0020275787,0.00072457193,0.00033968152,0.0018176979,0.01369324,0.016146885],"genre_scores_gemma":[0.94042164,0.0008188859,0.049759008,0.00034890816,0.000071803435,0.00014900893,0.0029915504,0.00024041405,0.0051988233],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992132,0.00037650074,0.000042740678,0.00019116246,0.00009591694,0.00008041653],"domain_scores_gemma":[0.9983329,0.0010820462,0.00007565868,0.00014592259,0.0002826534,0.000080800884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014587187,0.0016438946,0.00052362046,0.00070155767,0.00034540234,0.0012151488,0.0008573144,0.000718525,0.0024263312],"category_scores_gemma":[0.0046273717,0.00030544066,0.000722364,0.00031898855,0.0003218352,0.0019813823,0.00092946156,0.0017159685,0.00136668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016352207,0.00064138044,0.007581026,0.00031452306,0.00031815958,0.00029888668,0.0005891338,0.3015839,0.01819319,0.0049072113,0.010542265,0.6533951],"study_design_scores_gemma":[0.00002603079,0.00021788146,0.0009567266,0.00003255047,0.00006454281,0.000078862,0.0001317633,0.98798984,0.006850644,0.0018777801,0.0017448495,0.000028432862],"about_ca_topic_score_codex":0.011235763,"about_ca_topic_score_gemma":0.011434395,"teacher_disagreement_score":0.011235763,"about_ca_system_score_codex":0.00075361674,"about_ca_system_score_gemma":0.0010454028,"threshold_uncertainty_score":0.022340775},"labels":[],"label_agreement":null},{"id":"W33598619","doi":"","title":"Improving a statistical language model by modulating the effects of context words","year":2008,"lang":"en","type":"article","venue":"UCL Discovery (University College London)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Feature (linguistics); Word (group theory); Artificial intelligence; Context (archaeology); Feature vector; Language model; Artificial neural network; Context model; Natural language processing; Pattern recognition (psychology); Mathematics; Linguistics","score_opus":0.006199984077451072,"score_gpt":0.18540975469747722,"score_spread":0.17920977062002616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W33598619","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13207498,0.0012859855,0.85877913,0.0008385378,0.00027635775,0.000056809804,0.0002563239,0.0029567666,0.003475079],"genre_scores_gemma":[0.7803099,0.0009232033,0.21283253,0.0004663666,0.00028605794,0.00012311238,0.0007111902,0.0004677396,0.0038799653],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925023,0.00025897424,0.000041996576,0.00021341722,0.00016366277,0.0000715719],"domain_scores_gemma":[0.9971246,0.0019723645,0.00014035341,0.00038445232,0.0003148257,0.00006336168],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016954296,0.0013335465,0.0008780592,0.00058953825,0.00038982512,0.0011735315,0.0010771526,0.0009581466,0.0020748067],"category_scores_gemma":[0.007724555,0.0004814824,0.00083827606,0.00064835965,0.0004548528,0.0030483638,0.00084247655,0.0018041657,0.0016375943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000446879,0.00037512477,0.005049266,0.00019119843,0.00030924898,0.00015533416,0.00012941391,0.6139612,0.06006715,0.0060941656,0.0026561243,0.3105649],"study_design_scores_gemma":[0.000009307129,0.000054035867,0.0003237644,0.0000044858243,0.000038310074,0.00002149474,0.0000047010362,0.9919458,0.004788616,0.002395627,0.00040140134,0.000012401159],"about_ca_topic_score_codex":0.0049460633,"about_ca_topic_score_gemma":0.008151547,"teacher_disagreement_score":0.0049460633,"about_ca_system_score_codex":0.00039213896,"about_ca_system_score_gemma":0.0008130074,"threshold_uncertainty_score":0.009834528},"labels":[],"label_agreement":null},{"id":"W37346518","doi":"10.7717/peerj-cs.1363","title":"A Simple Closed-Class/Open-Class Factorization for Improved Language Modeling.","year":2001,"lang":"en","type":"article","venue":"NLPRS","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Zemkopības ministrija","keywords":"Perplexity; Language model; Class (philosophy); Computer science; Smoothing; Generalization; Function (biology); Simple (philosophy); Artificial intelligence; Natural language processing; Mathematics","score_opus":0.036167706508397304,"score_gpt":0.29543466461596823,"score_spread":0.25926695810757094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W37346518","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002825585,0.000111279776,0.98961276,0.00020897298,0.0001587875,0.000223159,0.00075034593,0.0032113453,0.0028978335],"genre_scores_gemma":[0.09122168,0.00017137123,0.8902231,0.0002436185,0.00016107314,0.000666845,0.0044092224,0.0010647792,0.01183824],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972988,0.0006348408,0.00024868408,0.00081279775,0.0006991182,0.0003056965],"domain_scores_gemma":[0.9970432,0.0012026898,0.00017196301,0.00060426706,0.00083774875,0.00014017956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029905674,0.0012323809,0.00080995803,0.001532922,0.0014022994,0.002692837,0.0023512165,0.0016744966,0.024992773],"category_scores_gemma":[0.009523786,0.0007657891,0.0031459418,0.0011347304,0.0008467557,0.004625705,0.0024165853,0.002932444,0.01315085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005385297,0.00044566236,0.0027605689,0.00039317002,0.00013363459,0.00043323304,0.00064574246,0.06156517,0.012632009,0.1175097,0.031315073,0.7716276],"study_design_scores_gemma":[0.000050265906,0.0000917454,0.0005668219,0.00006925564,0.00004352451,0.00026549413,0.00025494923,0.8593537,0.004988818,0.08811453,0.04614744,0.0000534578],"about_ca_topic_score_codex":0.012103994,"about_ca_topic_score_gemma":0.01210599,"teacher_disagreement_score":0.024992773,"about_ca_system_score_codex":0.0010789818,"about_ca_system_score_gemma":0.0031561626,"threshold_uncertainty_score":0.083609104},"labels":[],"label_agreement":null},{"id":"W38777760","doi":"10.1248/cpb.c24-00192","title":"Mots composés dans les modèles de langue pour la recherche d'information","year":2004,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Humanities; Philosophy","score_opus":0.18906402945056588,"score_gpt":0.32224133636519425,"score_spread":0.13317730691462837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W38777760","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015194104,0.0020388053,0.9551977,0.003102217,0.00042157833,0.000111732064,0.0021076926,0.0021178548,0.019708341],"genre_scores_gemma":[0.41093585,0.0050756866,0.5192247,0.0007181427,0.00042800262,0.001018595,0.00394412,0.0008384017,0.057816572],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998958,0.00043147014,0.00004832395,0.00019931493,0.0003079047,0.000054877662],"domain_scores_gemma":[0.9991321,0.0005012391,0.000075750155,0.0001365447,0.000114506256,0.000039903673],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014709871,0.0010381246,0.0009480674,0.0010411295,0.00070095324,0.0030515115,0.0009146285,0.0015942642,0.010811083],"category_scores_gemma":[0.0042702635,0.0004892582,0.0014027241,0.0013069498,0.0011934254,0.0033215613,0.0010981513,0.0018190034,0.0031775197],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002191335,0.00006579881,0.0015163647,0.0004457237,0.00015950589,0.000508385,0.0005673533,0.29089642,0.0056998245,0.597593,0.011075116,0.091253385],"study_design_scores_gemma":[0.00004233731,0.00004022267,0.0003749932,0.000083144114,0.000047573798,0.00013799005,0.00013628724,0.8016073,0.002902764,0.1413079,0.05328977,0.000029739085],"about_ca_topic_score_codex":0.013634706,"about_ca_topic_score_gemma":0.013997986,"teacher_disagreement_score":0.013634706,"about_ca_system_score_codex":0.0017751828,"about_ca_system_score_gemma":0.0016427853,"threshold_uncertainty_score":0.036166728},"labels":[],"label_agreement":null},{"id":"W40989923","doi":"10.1007/3-540-45108-0_71","title":"A Cognitive Model for Automatic Narrative Summarization in a Self-Educational System","year":2000,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Automatic summarization; Computer science; Reuse; Narrative; Field (mathematics); Generator (circuit theory); Artificial intelligence; Linguistics; Engineering","score_opus":0.021239410143157268,"score_gpt":0.2613588401457524,"score_spread":0.2401194300025951,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W40989923","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026407111,0.00010343407,0.96210945,0.00036965453,0.000033407,0.00024054969,0.00040616072,0.0042943927,0.0060358983],"genre_scores_gemma":[0.44946286,0.00015667829,0.5417075,0.00015782974,0.000047186957,0.00040209544,0.0010957144,0.00031150493,0.006658529],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934894,0.00018704474,0.00006512891,0.00019497825,0.00014700575,0.000056851968],"domain_scores_gemma":[0.9980446,0.0009948724,0.00014543664,0.0002399137,0.0004600541,0.0001151225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010606237,0.0005772972,0.0003851044,0.0013962174,0.000768223,0.0029554584,0.0017181752,0.0009903102,0.008749411],"category_scores_gemma":[0.0047147586,0.00047093528,0.001151896,0.0007625773,0.0007321049,0.0035370642,0.0011669981,0.00084463955,0.0016331708],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009250094,0.0007165063,0.007236409,0.0011155072,0.0003687613,0.00081044785,0.010174794,0.1154037,0.05219595,0.24219751,0.018793827,0.5500615],"study_design_scores_gemma":[0.000056992798,0.00010864105,0.0012828499,0.000057108628,0.00016081744,0.00018089055,0.000621782,0.9162011,0.009235682,0.06094356,0.011097382,0.000053213178],"about_ca_topic_score_codex":0.008944902,"about_ca_topic_score_gemma":0.008267843,"teacher_disagreement_score":0.008944902,"about_ca_system_score_codex":0.0010736703,"about_ca_system_score_gemma":0.0011184504,"threshold_uncertainty_score":0.029269755},"labels":[],"label_agreement":null},{"id":"W4135815","doi":"","title":"Improving the precision of a closed-domain question-answering system with semantic information","year":2004,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Concordia University","funders":"","keywords":"Information retrieval; Question answering; Computer science; Relevance (law); Hierarchy; Domain (mathematical analysis); Set (abstract data type); Information system; Semantics (computer science); Mathematics; Engineering","score_opus":0.0064219763785140944,"score_gpt":0.20381151127181837,"score_spread":0.19738953489330427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4135815","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5413767,0.0043485966,0.4068596,0.0025228313,0.0004388161,0.00094305445,0.0011274404,0.032007173,0.010375782],"genre_scores_gemma":[0.7581972,0.00062891195,0.23500459,0.0005663531,0.00028107042,0.00025845767,0.0023531357,0.0006126091,0.00209768],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98060143,0.009267387,0.0015046672,0.0032624702,0.0045998357,0.00076414435],"domain_scores_gemma":[0.85426176,0.122506835,0.0022522258,0.010569326,0.00972844,0.0006813642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021163464,0.0017670922,0.0031612075,0.0024849158,0.0020007114,0.0041953893,0.0032367753,0.005100796,0.0031845716],"category_scores_gemma":[0.13130865,0.0009172462,0.0012161105,0.0024056768,0.0015384465,0.008959996,0.0031074274,0.0033767393,0.0027181467],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0071556717,0.0021091343,0.01934489,0.0029692291,0.0008232099,0.00066088256,0.00496861,0.07070937,0.11031595,0.0046326984,0.0141872205,0.7621232],"study_design_scores_gemma":[0.001045101,0.0026761538,0.020780727,0.00025038983,0.0014381608,0.0011962183,0.0012596848,0.7301481,0.20843805,0.011154837,0.021170033,0.00044251134],"about_ca_topic_score_codex":0.008496326,"about_ca_topic_score_gemma":0.0035473865,"teacher_disagreement_score":0.021163464,"about_ca_system_score_codex":0.001818581,"about_ca_system_score_gemma":0.0014558038,"threshold_uncertainty_score":0.11192441},"labels":[],"label_agreement":null},{"id":"W41647797","doi":"10.1007/978-3-319-10061-6_14","title":"Answering Yes/No Questions in Legal Bar Exams","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Question answering; Sentence; Artificial intelligence; Natural language processing; Task (project management); Negation; Construct (python library); Knowledge base; Variety (cybernetics); Natural language; Domain (mathematical analysis); Programming language","score_opus":0.016528756874145904,"score_gpt":0.2411456648454812,"score_spread":0.2246169079713353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W41647797","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48221344,0.0027169678,0.12378997,0.02021088,0.002510763,0.0010084551,0.0038643624,0.0040395567,0.35964563],"genre_scores_gemma":[0.8593863,0.0007907717,0.054451153,0.0024524515,0.000788016,0.00038219368,0.004715375,0.0005251567,0.076508656],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9942673,0.003053308,0.00031713003,0.00067281025,0.0011431748,0.0005462132],"domain_scores_gemma":[0.9780742,0.017128743,0.0011859593,0.0007223326,0.0018685955,0.0010201908],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00542865,0.00057568285,0.0006538981,0.00089638267,0.0014503746,0.0035186121,0.0011434937,0.0021461148,0.03483853],"category_scores_gemma":[0.046410188,0.00043015997,0.0005510755,0.000686685,0.0007860018,0.0058331955,0.0025167488,0.00227912,0.011997291],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012223095,0.0017698177,0.053805914,0.000611019,0.00007324704,0.00066976465,0.010746297,0.0033924482,0.0054688742,0.04741181,0.18861297,0.6862156],"study_design_scores_gemma":[0.00026351766,0.0009946261,0.09658445,0.0020669925,0.00021932511,0.0013270311,0.026457101,0.048803754,0.016268432,0.2937201,0.5130827,0.0002120101],"about_ca_topic_score_codex":0.0015060869,"about_ca_topic_score_gemma":0.0024489288,"teacher_disagreement_score":0.03483853,"about_ca_system_score_codex":0.0007944166,"about_ca_system_score_gemma":0.0010978945,"threshold_uncertainty_score":0.11654651},"labels":[],"label_agreement":null},{"id":"W4200001833","doi":"10.1037/amp0000863","title":"Quantifying the selective forgetting and integration of ideas in science and technology.","year":2021,"lang":"en","type":"article","venue":"American Psychologist","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kellogg's (Canada)","funders":"Air Force Office of Scientific Research","keywords":"Forgetting; Collective memory; Retrieval-induced forgetting; Cognitive psychology; Process (computing); Cognitive science; Trademark; Psychology; Object (grammar); History; Sociology; Computer science; Political science; Artificial intelligence; Law","score_opus":0.03735793854743633,"score_gpt":0.34738091761382,"score_spread":0.31002297906638365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200001833","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91502213,0.007516809,0.064767405,0.0015146183,0.00009116459,0.00012508596,0.0010540922,0.0003012392,0.009607505],"genre_scores_gemma":[0.9881811,0.0007365487,0.009827189,0.00007772878,0.000048701168,0.000061923645,0.00041012117,0.000035597488,0.00062113476],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99382097,0.0018432131,0.0004937676,0.0011950253,0.0022866684,0.00036033388],"domain_scores_gemma":[0.86897224,0.08063354,0.026126228,0.013966672,0.0073957113,0.002905607],"candidate_categories":["bibliometrics","sts"],"consensus_categories":[],"category_scores_codex":[0.0118985,0.0005668741,0.0007645226,0.0086413175,0.0012022451,0.0032100622,0.0010724498,0.0014654374,0.0019327777],"category_scores_gemma":[0.12868325,0.00055977603,0.0012541385,0.0071192677,0.0025061723,0.009161451,0.0038890643,0.0017075911,0.00036900854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005818177,0.0003240333,0.5833805,0.0014654704,0.0012891277,0.000606344,0.011659352,0.028885452,0.006880558,0.034367148,0.002546518,0.32801372],"study_design_scores_gemma":[0.00004755677,0.000906674,0.66625816,0.00038601464,0.00070123974,0.0014120585,0.005812625,0.10037059,0.00936766,0.19897118,0.015541847,0.00022430393],"about_ca_topic_score_codex":0.004710907,"about_ca_topic_score_gemma":0.0041525904,"teacher_disagreement_score":0.9987978,"about_ca_system_score_codex":0.0021326826,"about_ca_system_score_gemma":0.0010428479,"threshold_uncertainty_score":0.062926054},"labels":[],"label_agreement":null},{"id":"W4200099029","doi":"10.22148/001c.30704","title":"Annotation Guidelines for Narrative Levels","year":2021,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Narrative; CLARITY; Rule of thumb; Focus (optics); Set (abstract data type); Computer science; Sentence; Narrative criticism; Linguistics; Narrative history; Natural language processing; Philosophy; Programming language","score_opus":0.21613679332068877,"score_gpt":0.39836663804300365,"score_spread":0.18222984472231488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200099029","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004316503,0.00069550116,0.8556529,0.005185537,0.0017178388,0.0030113794,0.029258104,0.020767648,0.07939458],"genre_scores_gemma":[0.021106089,0.0006005858,0.9048599,0.0016347115,0.00031854145,0.0054893116,0.026837155,0.010209744,0.028944021],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9888975,0.004857085,0.0026529576,0.0012577726,0.0018752564,0.00045954698],"domain_scores_gemma":[0.9481453,0.018135766,0.001863034,0.006317526,0.024541609,0.0009968068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010561371,0.0015857447,0.00070613436,0.0054134484,0.0033732266,0.0052143894,0.0025484615,0.0032347373,0.06105232],"category_scores_gemma":[0.051358078,0.0017144359,0.0010714026,0.003346609,0.0021692428,0.0075333924,0.004765513,0.0040862877,0.0426774],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034143683,0.0001266553,0.0014415307,0.0026724432,0.0000347355,0.0005915522,0.021994164,0.0011214844,0.01345447,0.2086916,0.56790566,0.18162431],"study_design_scores_gemma":[0.000023214752,0.000024112578,0.0005092261,0.0008945873,0.000020789672,0.00032337496,0.002023082,0.0024333748,0.0055452557,0.04540408,0.942726,0.00007280251],"about_ca_topic_score_codex":0.006419937,"about_ca_topic_score_gemma":0.009851269,"teacher_disagreement_score":0.06105232,"about_ca_system_score_codex":0.0023310126,"about_ca_system_score_gemma":0.0038552848,"threshold_uncertainty_score":0.20424038},"labels":[],"label_agreement":null},{"id":"W4200314022","doi":"10.1145/3487057","title":"Joined Type Length Encoding for Nested Named Entity Recognition","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Asian and Low-Resource Language Information Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Named-entity recognition; Encoding (memory); Leverage (statistics); Artificial intelligence; Sequence (biology); Pattern recognition (psychology); Natural language processing; Task (project management); Genetics","score_opus":0.018109367999372884,"score_gpt":0.24983490658388774,"score_spread":0.23172553858451486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200314022","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027949568,0.00070561393,0.9529697,0.00038022455,0.00037229687,0.00011056041,0.003713106,0.00939299,0.00440595],"genre_scores_gemma":[0.38518113,0.00061426824,0.5876678,0.0005548647,0.00016986187,0.0003116973,0.01587435,0.001047186,0.00857878],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99907243,0.00023329428,0.000146859,0.00028614252,0.00018653879,0.00007477348],"domain_scores_gemma":[0.99690956,0.0009438737,0.00025219895,0.0011654947,0.0006500544,0.00007887022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011306454,0.0007697615,0.0005367663,0.0012778364,0.00038364055,0.00090205553,0.00130523,0.00091230596,0.004609697],"category_scores_gemma":[0.00605398,0.0002723511,0.00065677334,0.0013993077,0.00055895495,0.005321773,0.0013617533,0.0013366401,0.00329912],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064731686,0.00014699483,0.0061008716,0.0005161803,0.0000982777,0.00052137405,0.00053094584,0.05363701,0.03197392,0.043069478,0.021815035,0.8409426],"study_design_scores_gemma":[0.00006163494,0.00030139266,0.0031623777,0.00020552507,0.00013256921,0.0008205065,0.00028270157,0.7432603,0.07871496,0.09117325,0.08169733,0.00018745637],"about_ca_topic_score_codex":0.0018988965,"about_ca_topic_score_gemma":0.0036632437,"teacher_disagreement_score":0.004609697,"about_ca_system_score_codex":0.0005532678,"about_ca_system_score_gemma":0.000830068,"threshold_uncertainty_score":0.015420973},"labels":[],"label_agreement":null},{"id":"W4200327391","doi":"10.22148/001c.30698","title":"Narrative Boundaries Annotation Guide","year":2021,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narrative; Scripting language; Annotation; Computer science; Schema (genetic algorithms); Narrative criticism; Narrative structure; World Wide Web; Information retrieval; Narrative inquiry; Artificial intelligence; Linguistics; Programming language; Philosophy","score_opus":0.03377907282816144,"score_gpt":0.3051666231684335,"score_spread":0.27138755034027207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200327391","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006752892,0.0010607733,0.39662808,0.0043554795,0.0020649044,0.0039691576,0.25289008,0.08031484,0.25196382],"genre_scores_gemma":[0.025961211,0.00094703684,0.5684124,0.0022486236,0.00033067647,0.007708895,0.19529717,0.031062383,0.16803157],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983236,0.0005729891,0.0003202117,0.00031452024,0.00035950425,0.00010910552],"domain_scores_gemma":[0.9886085,0.0046749013,0.0004630742,0.0014903507,0.004410656,0.00035252888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002965464,0.0012754829,0.0005796498,0.0049031903,0.0020062374,0.003243066,0.0014831998,0.001992797,0.19862106],"category_scores_gemma":[0.018313505,0.0009805528,0.00062940246,0.0028094652,0.0007296677,0.0036675155,0.00311691,0.0023990762,0.0950919],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000189877,0.000052426793,0.00066519016,0.0014655497,0.000011099594,0.00032882197,0.006374961,0.00071561034,0.005224338,0.0315942,0.8282213,0.12515667],"study_design_scores_gemma":[0.000009424624,0.0000057036577,0.00027604535,0.00020759123,0.0000036610436,0.00007522301,0.00046798144,0.0007830655,0.0014864748,0.0037097654,0.99295723,0.000017853507],"about_ca_topic_score_codex":0.0074088746,"about_ca_topic_score_gemma":0.013196697,"teacher_disagreement_score":0.19862106,"about_ca_system_score_codex":0.0016173304,"about_ca_system_score_gemma":0.002801918,"threshold_uncertainty_score":0.6644536},"labels":[],"label_agreement":null},{"id":"W4200445293","doi":"10.2196/27386","title":"Benchmarking Effectiveness and Efficiency of Deep Learning Models for Semantic Textual Similarity in the Clinical Domain: Validation Study","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Benchmarking; Computer science; Natural language processing; Similarity (geometry); Artificial intelligence; Deep learning; Domain (mathematical analysis); Semantic similarity; Information retrieval","score_opus":0.047565996480232776,"score_gpt":0.3516154646810429,"score_spread":0.3040494682008101,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200445293","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9463187,0.0069364766,0.030094331,0.0014575833,0.00046552694,0.0004582043,0.0045956555,0.005183927,0.004489583],"genre_scores_gemma":[0.9585772,0.000787773,0.027303608,0.00037642595,0.00010384209,0.00020637507,0.010862904,0.0002859054,0.0014959674],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9926031,0.00340779,0.0008930522,0.0016518576,0.00089686393,0.0005472776],"domain_scores_gemma":[0.9752363,0.016819334,0.0010993821,0.0022684599,0.0038012783,0.00077529764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013559267,0.003602884,0.0012874833,0.002459366,0.0007607902,0.00151998,0.0029516795,0.0028679888,0.0018002952],"category_scores_gemma":[0.029818391,0.00074568106,0.001590439,0.0017970368,0.0012312733,0.0024837672,0.002297773,0.002848074,0.0011881058],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004464249,0.0029220604,0.06129748,0.0017205109,0.0018464759,0.00039944882,0.00038070988,0.6700041,0.0059426017,0.001273036,0.022351047,0.22739823],"study_design_scores_gemma":[0.00020049386,0.0008858732,0.006110363,0.000083115825,0.00018639881,0.000081889535,0.00012346297,0.98428404,0.0060708425,0.0008938035,0.0010443231,0.00003540662],"about_ca_topic_score_codex":0.028969306,"about_ca_topic_score_gemma":0.023635356,"teacher_disagreement_score":0.028969306,"about_ca_system_score_codex":0.0033420236,"about_ca_system_score_gemma":0.0025983923,"threshold_uncertainty_score":0.0717091},"labels":[],"label_agreement":null},{"id":"W4200584900","doi":"10.1145/3426971","title":"SANTM: Efficient Self-attention-driven Network for Text Matching","year":2021,"lang":"en","type":"article","venue":"ACM Transactions on Internet Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Business Finland; Academy of Finland","keywords":"Computer science; Artificial intelligence; Leverage (statistics); Matching (statistics); Pooling; Inference; Question answering; Natural language processing; Paraphrase; Task (project management); Sentence; Benchmark (surveying); Machine learning","score_opus":0.013230423479802371,"score_gpt":0.2498152077177358,"score_spread":0.23658478423793342,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200584900","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045224555,0.0013778039,0.9312022,0.00090933323,0.00029268826,0.00037474456,0.00247793,0.012506178,0.0056346245],"genre_scores_gemma":[0.5635204,0.0008133158,0.40206268,0.0013082143,0.00035433783,0.0007695701,0.009297035,0.000782723,0.021091748],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995597,0.00009045187,0.000025164354,0.0001754633,0.00009597116,0.000053338976],"domain_scores_gemma":[0.9992167,0.00036468104,0.00007059808,0.00014122597,0.00015913422,0.00004761772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009935498,0.0009776601,0.0010055744,0.0014845429,0.00057120575,0.0007708641,0.0032156417,0.0020503718,0.0049569383],"category_scores_gemma":[0.003951554,0.0006118585,0.00097131607,0.0012677641,0.0005676911,0.0027598054,0.0016625293,0.0015788279,0.0020497437],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005653247,0.00046681808,0.0025499254,0.00038129423,0.00017968549,0.00028278716,0.0002692557,0.33173463,0.017100295,0.021185182,0.03703572,0.5882491],"study_design_scores_gemma":[0.00001562918,0.000026486827,0.00012958003,0.000005335691,0.000012970114,0.000028855671,0.000009485193,0.9891382,0.0019844673,0.007152733,0.001490332,0.0000058165915],"about_ca_topic_score_codex":0.0076509337,"about_ca_topic_score_gemma":0.013050331,"teacher_disagreement_score":0.0076509337,"about_ca_system_score_codex":0.0017175882,"about_ca_system_score_gemma":0.0013004866,"threshold_uncertainty_score":0.016582608},"labels":[],"label_agreement":null},{"id":"W4200631513","doi":"10.1609/aaai.v36i10.21325","title":"Predicting Above-Sentence Discourse Structure Using Distant Supervision from Topic Segmentation","year":2022,"lang":"en","type":"preprint","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Natural language processing; Computer science; Paragraph; Sentence; Parsing; Artificial intelligence; Automatic summarization; Segmentation; Task (project management); Semantic role labeling; Linguistics","score_opus":0.09739871542539401,"score_gpt":0.33011612829842624,"score_spread":0.23271741287303221,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200631513","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48323777,0.004369478,0.48723307,0.0017210052,0.00038116841,0.00028424588,0.0062359544,0.007350455,0.0091868555],"genre_scores_gemma":[0.8655366,0.0006526701,0.1172544,0.00015289089,0.00036697977,0.0001811691,0.011066326,0.000326082,0.004462893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993451,0.00024920554,0.00003489426,0.00025185445,0.0000614904,0.00005740813],"domain_scores_gemma":[0.99600357,0.002631008,0.00037281855,0.00024367236,0.0005821594,0.00016673587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012608058,0.0010064533,0.0005593376,0.0019909532,0.0007728003,0.001081976,0.0006702397,0.0011344381,0.002449255],"category_scores_gemma":[0.004931436,0.0003922715,0.0006789248,0.0011603704,0.00038902546,0.0021244476,0.0008883677,0.0014640411,0.0021963008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021147283,0.0008138742,0.038721118,0.00118238,0.00030158524,0.0010034565,0.0048697162,0.06647521,0.09946089,0.009764893,0.039628394,0.7356638],"study_design_scores_gemma":[0.000087015185,0.00021965067,0.015340385,0.00010345852,0.00016172256,0.00018571687,0.0006856148,0.92850614,0.026678579,0.015581858,0.0123991985,0.000050743023],"about_ca_topic_score_codex":0.0037733498,"about_ca_topic_score_gemma":0.008955548,"teacher_disagreement_score":0.0037733498,"about_ca_system_score_codex":0.00059062004,"about_ca_system_score_gemma":0.0009635864,"threshold_uncertainty_score":0.0081935525},"labels":[],"label_agreement":null},{"id":"W4200633215","doi":"10.1609/aaai.v36i10.21293","title":"From Good to Best: Two-Stage Training for Cross-Lingual Machine Reading Comprehension","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Set (abstract data type); Natural language processing; Reading (process); Comprehension; Task (project management); Resource (disambiguation); Reading comprehension; Machine learning; Recall; Precision and recall; Linguistics","score_opus":0.172198085721702,"score_gpt":0.3666931078557866,"score_spread":0.19449502213408462,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200633215","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31117898,0.0032467418,0.643713,0.0027078628,0.00043954354,0.001013821,0.0015954344,0.026678981,0.009425698],"genre_scores_gemma":[0.7401291,0.00023654172,0.2467435,0.0015985481,0.00017151628,0.00076316233,0.0044452376,0.00075615215,0.0051562143],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960335,0.0013683754,0.0002546047,0.0015017196,0.0004567903,0.0003851183],"domain_scores_gemma":[0.98595285,0.009406245,0.00041581952,0.0016736413,0.0020011198,0.00055030425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008368791,0.0033921266,0.0021620183,0.002058301,0.0012136376,0.0024514638,0.005534811,0.004948859,0.0062203137],"category_scores_gemma":[0.02096722,0.0015742539,0.0017834082,0.0012413568,0.001371187,0.005462499,0.004925737,0.006330049,0.003291523],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025920123,0.0020842187,0.021247933,0.000613252,0.0005915003,0.00068767084,0.0012621536,0.194021,0.016424121,0.004473211,0.023245059,0.73275775],"study_design_scores_gemma":[0.0000949644,0.00030394405,0.0011303338,0.000035652258,0.00006458161,0.00007576365,0.00010707645,0.9868129,0.004878275,0.0054868986,0.0009710355,0.00003874991],"about_ca_topic_score_codex":0.005241988,"about_ca_topic_score_gemma":0.008757806,"teacher_disagreement_score":0.008368791,"about_ca_system_score_codex":0.0014354562,"about_ca_system_score_gemma":0.0019994979,"threshold_uncertainty_score":0.044258893},"labels":[],"label_agreement":null},{"id":"W4200635123","doi":"10.18653/v1/2022.naacl-main.168","title":"GPL: Generative Pseudo Labeling for Unsupervised Domain Adaptation of Dense Retrieval","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Deutsche Forschungsgemeinschaft","keywords":"Generative grammar; Computer science; Domain (mathematical analysis); Natural language processing; Artificial intelligence; Adaptation (eye); Domain adaptation; Mathematics; Psychology; Neuroscience","score_opus":0.024539071504083444,"score_gpt":0.25158812013711035,"score_spread":0.2270490486330269,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200635123","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015698022,0.00026285753,0.974918,0.00011773343,0.00007475739,0.00007363468,0.00063330814,0.02133956,0.0010103625],"genre_scores_gemma":[0.07191544,0.00037851653,0.9042158,0.0006378339,0.00017126038,0.0004557137,0.008243828,0.0062438645,0.007737737],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99845076,0.0005295406,0.00006834036,0.00045338282,0.00034487812,0.00015310473],"domain_scores_gemma":[0.99762625,0.0008956933,0.000087759705,0.0009163746,0.0003576635,0.000116268035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022299953,0.0016135505,0.0018613932,0.0021697052,0.0012086441,0.002236912,0.0048288805,0.0028155674,0.01605406],"category_scores_gemma":[0.006543548,0.0015898023,0.0020242175,0.0026280081,0.0013302161,0.0040063146,0.0050769653,0.0032193977,0.015135952],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000510035,0.00026005707,0.0008468848,0.00036920345,0.00019633677,0.00022771423,0.00031167973,0.11139795,0.011219979,0.024251122,0.10982287,0.7405862],"study_design_scores_gemma":[0.000077278986,0.000045172917,0.00015825497,0.000021534019,0.000025686206,0.00009861207,0.000046952162,0.94915754,0.0046111504,0.034484223,0.011239455,0.000034106863],"about_ca_topic_score_codex":0.0119986925,"about_ca_topic_score_gemma":0.024200583,"teacher_disagreement_score":0.01605406,"about_ca_system_score_codex":0.0013767076,"about_ca_system_score_gemma":0.0018732535,"threshold_uncertainty_score":0.05370617},"labels":[],"label_agreement":null},{"id":"W4205226497","doi":"10.1109/smc52423.2021.9658712","title":"Is Timing Critical to Trace Reconstruction?","year":2021,"lang":"en","type":"article","venue":"2021 IEEE International Conference on Systems, Man, and Cybernetics (SMC)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"TRACE (psycholinguistics); Computer science","score_opus":0.09793658350581186,"score_gpt":0.33342699258391295,"score_spread":0.2354904090781011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205226497","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09223968,0.0036119104,0.8910524,0.005523317,0.00050183746,0.000057382054,0.00062709564,0.0028968956,0.0034894538],"genre_scores_gemma":[0.93033075,0.0021798336,0.06377714,0.0005083356,0.0003479169,0.000063817686,0.0006819608,0.000515652,0.0015945281],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99819285,0.00050009834,0.00014481047,0.00054741086,0.00037683398,0.0002379851],"domain_scores_gemma":[0.9815732,0.011502262,0.002491349,0.0020842855,0.001703712,0.00064520235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031959128,0.0010198244,0.0011257452,0.0012705748,0.0007351299,0.0020158084,0.001929438,0.0016395366,0.0022128283],"category_scores_gemma":[0.038601443,0.0006859776,0.0007276845,0.001291037,0.0016903888,0.0057008453,0.001369094,0.0027173415,0.0007494004],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006952825,0.0001891058,0.049747143,0.00076201506,0.00019336116,0.00047579117,0.0006497975,0.5840498,0.009905857,0.062035464,0.00779091,0.2835054],"study_design_scores_gemma":[0.000027318683,0.00006925295,0.0032564832,0.00009637127,0.000058866433,0.00025998466,0.00025348642,0.90912616,0.004189562,0.078639165,0.0039825896,0.000040791416],"about_ca_topic_score_codex":0.014505774,"about_ca_topic_score_gemma":0.010539429,"teacher_disagreement_score":0.014505774,"about_ca_system_score_codex":0.0011741426,"about_ca_system_score_gemma":0.002810884,"threshold_uncertainty_score":0.028842688},"labels":[],"label_agreement":null},{"id":"W4205635927","doi":"10.1162/coli_a_00426","title":"Deep Learning for Text Style Transfer: A Survey","year":2021,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":162,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Politeness; Style (visual arts); Task (project management); Variety (cybernetics); Natural language processing; Artificial intelligence; Transfer of learning; Field (mathematics); Natural language; Deep learning; Natural (archaeology); Linguistics","score_opus":0.04084954962181789,"score_gpt":0.2874410351691519,"score_spread":0.24659148554733404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205635927","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05410707,0.28961685,0.618791,0.0056002056,0.0015020344,0.00030203265,0.002325285,0.0052427575,0.022512734],"genre_scores_gemma":[0.60210806,0.15075538,0.20760031,0.0022134446,0.0030035425,0.00049396016,0.0086853625,0.0008581356,0.024281783],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.999316,0.00020623788,0.0000646911,0.00020253338,0.00015652044,0.00005407097],"domain_scores_gemma":[0.9985286,0.00084870536,0.0000779972,0.00023252088,0.00024587876,0.00006627622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016144832,0.0012699026,0.0011317789,0.0016625505,0.00032817916,0.0013114737,0.0015718447,0.0009088408,0.0041676373],"category_scores_gemma":[0.0042387987,0.00044085854,0.00093427557,0.0020166123,0.0004133914,0.0022646291,0.0012877875,0.0016816445,0.0022172546],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009439478,0.0002133605,0.0020236133,0.00091770856,0.0001480327,0.000039285635,0.00009490253,0.027951201,0.0021046423,0.0070018326,0.016577754,0.94283336],"study_design_scores_gemma":[0.000057661593,0.00027809327,0.004819018,0.00084091816,0.00018482372,0.00017050563,0.00017938654,0.86500084,0.006584863,0.048879836,0.07294323,0.000060688377],"about_ca_topic_score_codex":0.0027905274,"about_ca_topic_score_gemma":0.0025604214,"teacher_disagreement_score":0.0041676373,"about_ca_system_score_codex":0.0006905592,"about_ca_system_score_gemma":0.00074449,"threshold_uncertainty_score":0.013942182},"labels":[],"label_agreement":null},{"id":"W4205795307","doi":"10.1109/bigdata52589.2021.9671523","title":"BIRD-QA: A BERT-based Information Retrieval Approach to Domain Specific Question Answering","year":2021,"lang":"en","type":"article","venue":"2021 IEEE International Conference on Big Data (Big Data)","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Question answering; Information retrieval; Context (archaeology); Knowledge base; Preprocessor; Domain (mathematical analysis); Matching (statistics); Task (project management); F1 score; Language model; Artificial intelligence","score_opus":0.28268995294210725,"score_gpt":0.33270897878101213,"score_spread":0.050019025838904885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205795307","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011135208,0.0015012933,0.9594685,0.0007624832,0.00011945901,0.00076406024,0.0027354835,0.019156797,0.0043567214],"genre_scores_gemma":[0.21619684,0.0009338277,0.76139474,0.00096644944,0.00014964957,0.00071098184,0.011307348,0.00044677214,0.007893301],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984232,0.00056898873,0.00013110525,0.00044799587,0.00033968643,0.00008892828],"domain_scores_gemma":[0.9973355,0.0013911307,0.00011949479,0.00040060398,0.0006305572,0.00012259973],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027043791,0.0013468163,0.0009498996,0.004221884,0.00069611985,0.0014222906,0.003460765,0.0017384737,0.0049375542],"category_scores_gemma":[0.006847292,0.00067043136,0.0017246463,0.002222123,0.0007071688,0.0037462655,0.0018099862,0.0019450606,0.003401426],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006184633,0.0007852894,0.006166344,0.001163487,0.00036355434,0.00042671472,0.0012790586,0.14663525,0.014334899,0.028958084,0.064578995,0.7346898],"study_design_scores_gemma":[0.00003696868,0.00023592437,0.0015809701,0.00006594667,0.000071384144,0.00023631055,0.00016332952,0.94102263,0.0043710284,0.019872002,0.032288317,0.000055117376],"about_ca_topic_score_codex":0.023486946,"about_ca_topic_score_gemma":0.024945617,"teacher_disagreement_score":0.023486946,"about_ca_system_score_codex":0.0015752631,"about_ca_system_score_gemma":0.0012440995,"threshold_uncertainty_score":0.046700478},"labels":[],"label_agreement":null},{"id":"W4206850841","doi":"10.18653/v1/2022.acl-long.236","title":"Hallucinated but Factual! Inspecting the Factuality of Hallucinations in Abstractive Summarization","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Samsung; Compute Canada; Canadian Institute for Advanced Research","keywords":"Hallucinating; Automatic summarization; Computer science; Artificial intelligence; Natural language processing; Reinforcement learning; Machine learning","score_opus":0.011617738111289745,"score_gpt":0.23785216499537204,"score_spread":0.2262344268840823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206850841","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13938233,0.0012885344,0.8502845,0.0010786596,0.00018654707,0.00015311294,0.00046221862,0.0044650673,0.002699061],"genre_scores_gemma":[0.8044072,0.00041346712,0.19214967,0.00019802352,0.00014792064,0.000053370855,0.00062299607,0.00021852311,0.00178892],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99915457,0.0003275198,0.00007585654,0.0001895581,0.00020040924,0.000051978703],"domain_scores_gemma":[0.99231416,0.0040770303,0.0012419176,0.0008491764,0.0012302028,0.0002874216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018098346,0.00082035136,0.00053142203,0.0009964112,0.00045041562,0.001718819,0.0008188377,0.0008952146,0.0019394286],"category_scores_gemma":[0.015182745,0.00035601886,0.00030372362,0.00048360997,0.0007147875,0.0031606262,0.0014676102,0.0012236334,0.000757066],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016174869,0.00017330518,0.012878303,0.0010617545,0.00020509063,0.001108756,0.003150715,0.044808008,0.1704586,0.017854318,0.012424942,0.7342587],"study_design_scores_gemma":[0.00010226326,0.0008298272,0.014373773,0.00018173529,0.00022773832,0.0011312988,0.0012081101,0.77843815,0.13187067,0.050507072,0.020971697,0.00015769107],"about_ca_topic_score_codex":0.0008900338,"about_ca_topic_score_gemma":0.0014993259,"teacher_disagreement_score":0.0019394286,"about_ca_system_score_codex":0.00040024307,"about_ca_system_score_gemma":0.0003853045,"threshold_uncertainty_score":0.009571433},"labels":[],"label_agreement":null},{"id":"W4210588147","doi":"10.1145/3511322.3511327","title":"AI education matters","year":2021,"lang":"en","type":"article","venue":"AI Matters","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Impact; University of Toronto","funders":"","keywords":"Autoencoder; Artificial intelligence; Computer science; Sentence; Sequence (biology); Deep learning; Ideal (ethics); Word (group theory); Artificial neural network; Noise reduction; Natural language processing; Linguistics","score_opus":0.011002925934065646,"score_gpt":0.25754602420791567,"score_spread":0.24654309827385001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210588147","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013873756,0.006716474,0.005405606,0.68236953,0.054978576,0.000058034788,0.0010250239,0.0016945594,0.24636489],"genre_scores_gemma":[0.045775287,0.009541098,0.005106451,0.20004602,0.034481667,0.0002006176,0.0018167425,0.002317794,0.70071435],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953967,0.0009293444,0.00021856505,0.0008212664,0.0018327333,0.0008014136],"domain_scores_gemma":[0.9809996,0.0033738425,0.00091664854,0.001569043,0.00486497,0.008275873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040021595,0.0010746226,0.0006399204,0.0014098281,0.0040786164,0.010610065,0.0015949059,0.0060223625,0.27803162],"category_scores_gemma":[0.028458001,0.00039895004,0.0007215521,0.0016845295,0.0027427804,0.011798361,0.0050939815,0.007458567,0.18880011],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009374,0.000023752946,0.00028601685,0.000062544284,0.0000027791293,0.00004049462,0.00016350723,0.000071700124,0.000074264324,0.020642893,0.93537086,0.043251835],"study_design_scores_gemma":[0.0000026136881,0.000009262636,0.0002063949,0.00009900368,0.0000019335105,0.00009088675,0.00026347607,0.00010510732,0.00005614241,0.016467843,0.98269224,0.0000050644285],"about_ca_topic_score_codex":0.002702168,"about_ca_topic_score_gemma":0.005639245,"teacher_disagreement_score":0.27803162,"about_ca_system_score_codex":0.004082616,"about_ca_system_score_gemma":0.00810293,"threshold_uncertainty_score":0.9301084},"labels":[],"label_agreement":null},{"id":"W4210778341","doi":"10.1145/3474555","title":"Generating Factoid Questions with Question Type Enhanced Representation and Attention-based Copy Mechanism","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Asian and Low-Resource Language Information Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Mechanism (biology); Representation (politics); Interrogative; Benchmark (surveying); Encoder; Quality (philosophy); Exploit; Artificial intelligence; Key (lock); Focus (optics); Mode (computer interface); Natural language processing; Linguistics; Human–computer interaction","score_opus":0.00789677647571175,"score_gpt":0.2444467531973456,"score_spread":0.23654997672163386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210778341","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04187302,0.00085410726,0.93568355,0.00084911817,0.00019561387,0.00045133627,0.0012382811,0.014339479,0.004515533],"genre_scores_gemma":[0.3499547,0.0005422931,0.63196343,0.00066137564,0.00016847019,0.00046666473,0.006057566,0.00087162416,0.009313925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987758,0.00043605757,0.000083698054,0.00038864888,0.0002419049,0.00007385989],"domain_scores_gemma":[0.99622524,0.002168415,0.00018626619,0.0007350556,0.00056978024,0.00011527595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016531212,0.0013484624,0.0009124598,0.0015880223,0.00046980538,0.0013239938,0.0017581827,0.0015936814,0.007825831],"category_scores_gemma":[0.0090687005,0.00044044337,0.0013452023,0.0009524807,0.00068039156,0.0043839095,0.0022334165,0.001905284,0.0026344927],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003672537,0.00029664428,0.0025752385,0.00066468614,0.00011384455,0.00037884252,0.0009665802,0.015132063,0.039756883,0.012712245,0.02034951,0.90668607],"study_design_scores_gemma":[0.00016413844,0.00037328765,0.0027179962,0.00009116985,0.00024071055,0.00095349795,0.00053473125,0.854339,0.07209875,0.040297605,0.028087385,0.0001016854],"about_ca_topic_score_codex":0.002415364,"about_ca_topic_score_gemma":0.0034511823,"teacher_disagreement_score":0.007825831,"about_ca_system_score_codex":0.00084385637,"about_ca_system_score_gemma":0.0011166784,"threshold_uncertainty_score":0.02617997},"labels":[],"label_agreement":null},{"id":"W4210825693","doi":"10.1038/s44159-022-00025-3","title":"Instance theory as a domain-general framework for cognitive psychology","year":2022,"lang":"en","type":"article","venue":"Nature Reviews Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":49,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge; McGill University; University of Manitoba","funders":"","keywords":"Semantic memory; Episodic memory; Cognitive science; Cognition; Reconstructive memory; Cognitive psychology; Perspective (graphical); Associative property; Psychology; Cognitive architecture; Domain specificity; Computer science; Artificial intelligence; Explicit memory","score_opus":0.039063685941395754,"score_gpt":0.4177315447928267,"score_spread":0.37866785885143095,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210825693","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0058150403,0.0071258615,0.9677692,0.008144285,0.00020039888,0.000063756845,0.00030338357,0.00044495176,0.010133039],"genre_scores_gemma":[0.5055047,0.013003231,0.46670768,0.0037929828,0.002739057,0.0007486497,0.0015928261,0.00046186708,0.00544906],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99598074,0.0021923035,0.00027862724,0.0007536666,0.00060703483,0.00018758368],"domain_scores_gemma":[0.9928266,0.0052771876,0.00034871482,0.0008361716,0.00038198556,0.0003294671],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057550054,0.0015030389,0.0024487448,0.0052253753,0.001333516,0.0077464962,0.0037330769,0.00333416,0.004878109],"category_scores_gemma":[0.009551267,0.00096180977,0.0038807713,0.004424703,0.006008125,0.014486851,0.0033235855,0.008441499,0.0010814067],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013877685,0.000015982536,0.0002819617,0.0001816134,0.00008104968,0.000035339126,0.00032272513,0.0025793412,0.00028014655,0.98302734,0.0013371463,0.011843478],"study_design_scores_gemma":[0.0000055395167,0.0000056719045,0.00015480784,0.000027645754,0.00001600018,0.000040398718,0.00003680998,0.007403311,0.00007043874,0.98917013,0.0030616315,0.0000075209996],"about_ca_topic_score_codex":0.0022746648,"about_ca_topic_score_gemma":0.0014219776,"teacher_disagreement_score":0.0077464962,"about_ca_system_score_codex":0.002525003,"about_ca_system_score_gemma":0.0016398084,"threshold_uncertainty_score":0.030435741},"labels":[],"label_agreement":null},{"id":"W4212984080","doi":"10.2196/preprints.27210","title":"A Question-and-Answer System to Extract Data From Free-Text Oncological Pathology Reports (CancerBERT Network): Development Study (Preprint)","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Alberta Health Services","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Terminology; Named-entity recognition; Information retrieval; Pathology; Medicine; Linguistics","score_opus":0.06753405355201075,"score_gpt":0.3209014352907348,"score_spread":0.25336738173872403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4212984080","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23945439,0.002021154,0.5069705,0.0025265703,0.0009191301,0.0029532234,0.021801969,0.2072304,0.016122593],"genre_scores_gemma":[0.30706888,0.0008357311,0.6000243,0.0011841009,0.00016555621,0.0015644022,0.056421902,0.0018353062,0.030899836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925166,0.00017898165,0.000054109998,0.00027746265,0.00018181893,0.000056007186],"domain_scores_gemma":[0.9968612,0.0017584957,0.00011335942,0.00028345382,0.0008183693,0.00016503154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021424794,0.00091926195,0.00048495538,0.0009022863,0.00034571497,0.0009225119,0.0016884739,0.0012885978,0.012041497],"category_scores_gemma":[0.005532905,0.00042045818,0.0006376435,0.0005414914,0.00031257322,0.0031311929,0.0011227166,0.0011141059,0.008474491],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013301865,0.0014720487,0.007265956,0.0012735524,0.00023307125,0.0010443457,0.0009420473,0.02126234,0.045104012,0.005732754,0.12700158,0.787338],"study_design_scores_gemma":[0.0003231461,0.0010474109,0.0070257294,0.00013308665,0.00015585529,0.00092604145,0.00056077284,0.83281773,0.082350306,0.0031883658,0.07135299,0.0001185602],"about_ca_topic_score_codex":0.008439544,"about_ca_topic_score_gemma":0.008209541,"teacher_disagreement_score":0.012041497,"about_ca_system_score_codex":0.0010352673,"about_ca_system_score_gemma":0.0015381405,"threshold_uncertainty_score":0.040282845},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"bench_or_experimental","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"medium"}],"label_agreement":"split"},{"id":"W4213191780","doi":"10.1007/s12626-022-00105-z","title":"Overview and Discussion of the Competition on Legal Information Extraction/Entailment (COLIEE) 2021","year":2022,"lang":"en","type":"article","venue":"The Review of Socionetwork Strategies","topic":"Topic Modeling","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Hokkaido University; Shizuoka University; University of Alberta; National Institute of Informatics; Alberta Machine Intelligence Institute","keywords":"Task (project management); Statute; Computer science; Logical consequence; Component (thermodynamics); Competition (biology); Natural language processing; Information retrieval; Law; Artificial intelligence; Political science","score_opus":0.020410539368644436,"score_gpt":0.28133437493635527,"score_spread":0.2609238355677108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213191780","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06291305,0.09021676,0.28323188,0.12089659,0.056547917,0.00946552,0.09895205,0.04033627,0.23743999],"genre_scores_gemma":[0.102721244,0.01731185,0.29382637,0.02121551,0.00951148,0.00655691,0.35854,0.021954654,0.1683621],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9531818,0.019192448,0.0028342656,0.004494868,0.016759919,0.0035367387],"domain_scores_gemma":[0.9140526,0.028346084,0.0015578545,0.007124057,0.039504204,0.009415219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.067475215,0.0032466687,0.0034431953,0.012307995,0.0058367695,0.011581243,0.0062766173,0.004500407,0.03971196],"category_scores_gemma":[0.068088435,0.0011075283,0.0027728213,0.012420965,0.002067359,0.010028316,0.009134989,0.00476075,0.026311887],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038129714,0.0004255842,0.00089667336,0.0012170629,0.00009700645,0.00008134999,0.0004456859,0.001738109,0.002032159,0.0054452936,0.8278228,0.15941714],"study_design_scores_gemma":[0.00017929499,0.0003106623,0.00396632,0.00051852234,0.00005700268,0.00019557307,0.0006071583,0.009491152,0.0042914697,0.006774875,0.973464,0.0001439378],"about_ca_topic_score_codex":0.028812494,"about_ca_topic_score_gemma":0.052695286,"teacher_disagreement_score":0.067475215,"about_ca_system_score_codex":0.008516013,"about_ca_system_score_gemma":0.011698908,"threshold_uncertainty_score":0.3568473},"labels":[],"label_agreement":null},{"id":"W4214722770","doi":"10.24124/2018/58991","title":"A graph-based approach towards automatic text summarization.","year":2018,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval; Graph; Text graph; Multi-document summarization; The Internet; Natural language processing; Text processing; Artificial intelligence; World Wide Web; Theoretical computer science","score_opus":0.020134889870311777,"score_gpt":0.2567149650776157,"score_spread":0.23658007520730392,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4214722770","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020125043,0.0010841354,0.9903414,0.00016335037,0.000081774204,0.00017570975,0.0006930603,0.00446149,0.0009864065],"genre_scores_gemma":[0.03833218,0.0009267576,0.95224637,0.0001222545,0.00014725319,0.00026765763,0.0038187087,0.0004301094,0.0037087118],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990977,0.00026920892,0.000066078646,0.00028439335,0.000242394,0.00004020218],"domain_scores_gemma":[0.9988385,0.0004894406,0.0001421989,0.0001717664,0.00031749645,0.00004076356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00083405763,0.0015309324,0.0008906438,0.004464305,0.00055437686,0.0013903102,0.0013172815,0.0009904447,0.003086336],"category_scores_gemma":[0.0028451993,0.00043079996,0.001568987,0.0030914447,0.00038518317,0.0015299625,0.00077938987,0.0011078613,0.0024600823],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014651667,0.00015721,0.00058690854,0.0010761073,0.00031201585,0.00022791109,0.0003489781,0.078655176,0.032049287,0.014654493,0.029497914,0.84228754],"study_design_scores_gemma":[0.00005809803,0.00023561006,0.0011719947,0.00007817154,0.00021071122,0.00034841703,0.00020456203,0.8805173,0.017362645,0.046538662,0.053211838,0.000061922525],"about_ca_topic_score_codex":0.00435169,"about_ca_topic_score_gemma":0.006950276,"teacher_disagreement_score":0.004464305,"about_ca_system_score_codex":0.0006747404,"about_ca_system_score_gemma":0.00078409904,"threshold_uncertainty_score":0.010324776},"labels":[],"label_agreement":null},{"id":"W4220732073","doi":"10.5121/csit.2022.120615","title":"An IR-based QA System for Impact of Social Determinants of Health on Covid-19","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"","keywords":"Computer science; Ontology; Semantics (computer science); Metric (unit); Natural language; Natural language processing; Question answering; Coronavirus disease 2019 (COVID-19); Information retrieval; Artificial intelligence; Data science; World Wide Web; Medicine","score_opus":0.09158800524228218,"score_gpt":0.4013768906598264,"score_spread":0.30978888541754424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220732073","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.086705394,0.0039603305,0.5126496,0.0068907733,0.0010311197,0.0040119686,0.13912337,0.21940155,0.02622591],"genre_scores_gemma":[0.24061923,0.001058828,0.60119224,0.0018250069,0.0003414291,0.0019689382,0.14359497,0.0025266395,0.006872665],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99722135,0.0009782452,0.0004325271,0.00078543887,0.0004726647,0.000109806744],"domain_scores_gemma":[0.9878897,0.0074574654,0.00059362216,0.0010958968,0.0026678971,0.00029540903],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071301768,0.0009559331,0.001113759,0.0074140946,0.0012014495,0.0020190226,0.0013763064,0.0015640202,0.009802899],"category_scores_gemma":[0.020524269,0.00042883857,0.0013617967,0.0038268045,0.00038328703,0.003191362,0.0023173862,0.0010761705,0.004409056],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015582324,0.00096551346,0.03263972,0.004839689,0.0004988457,0.0011790558,0.0043985825,0.011736448,0.04764025,0.014209974,0.31756818,0.5627655],"study_design_scores_gemma":[0.00087671675,0.00082428945,0.048349373,0.0008159326,0.0008484773,0.0014016288,0.0027306105,0.3828982,0.05971118,0.030591028,0.47049925,0.00045328983],"about_ca_topic_score_codex":0.0099790655,"about_ca_topic_score_gemma":0.008947421,"teacher_disagreement_score":0.0099790655,"about_ca_system_score_codex":0.0014655078,"about_ca_system_score_gemma":0.0023760532,"threshold_uncertainty_score":0.0377084},"labels":[],"label_agreement":null},{"id":"W4220732108","doi":"10.1162/coli_a_00434","title":"Domain Adaptation with Pre-trained Transformers for Query-Focused Abstractive Text Summarization","year":2022,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Automatic summarization; Computer science; Transformer; Domain adaptation; Natural language processing; Adaptation (eye); Artificial intelligence; Transfer of learning; Task (project management); Information retrieval","score_opus":0.01946222096216806,"score_gpt":0.24978835952530382,"score_spread":0.23032613856313575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220732108","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041613773,0.0010438336,0.940353,0.0003195229,0.00010375061,0.00019705274,0.00067497167,0.014135673,0.0015584305],"genre_scores_gemma":[0.60281426,0.00071378774,0.38248953,0.00040513487,0.0001898425,0.00034689924,0.008061384,0.0006782035,0.00430091],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990835,0.00037561485,0.00007757233,0.0002740491,0.00013569584,0.0000535664],"domain_scores_gemma":[0.9973345,0.0012865303,0.00021876594,0.0005303778,0.0005381665,0.00009171294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001941258,0.0009420176,0.0007263044,0.0012157938,0.0003217999,0.0008211022,0.001438743,0.0007898915,0.0017542887],"category_scores_gemma":[0.0074727787,0.0003071917,0.0007701556,0.0010222304,0.00041843313,0.0025344167,0.0012331288,0.0018037517,0.0019102714],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058766274,0.00047530895,0.0024215605,0.00049321406,0.00015331422,0.0002442715,0.0005649574,0.18029541,0.054235484,0.0056987167,0.018757802,0.7360723],"study_design_scores_gemma":[0.00006171205,0.00023279183,0.0006004894,0.000018870429,0.000055442444,0.00008781966,0.00012432103,0.96212,0.024383925,0.007388795,0.0049020047,0.000023909895],"about_ca_topic_score_codex":0.0017716949,"about_ca_topic_score_gemma":0.0028731562,"teacher_disagreement_score":0.001941258,"about_ca_system_score_codex":0.00062581024,"about_ca_system_score_gemma":0.00090411375,"threshold_uncertainty_score":0.010266483},"labels":[],"label_agreement":null},{"id":"W4220759704","doi":"10.3390/app12062891","title":"Survey of BERT-Base Models for Scientific Text Classification: COVID-19 Case Study","year":2022,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":86,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Coronavirus disease 2019 (COVID-19); Computer science; Scientific literature; Pandemic; Data science; Context (archaeology); Task (project management); Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Domain (mathematical analysis); Artificial intelligence; History; Infectious disease (medical specialty); Medicine; Engineering; Disease; Mathematics","score_opus":0.26733874868496865,"score_gpt":0.35752594764442186,"score_spread":0.09018719895945321,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220759704","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17154603,0.20914388,0.4808636,0.023201307,0.0038137676,0.0020320849,0.048353057,0.03210949,0.028936809],"genre_scores_gemma":[0.5421573,0.049751356,0.2730275,0.004327251,0.0028035918,0.0012184627,0.10823043,0.0014516498,0.017032443],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967399,0.001401195,0.0004449753,0.0005424219,0.00068774057,0.0001837761],"domain_scores_gemma":[0.9820239,0.013488947,0.0005702202,0.0010063952,0.0023303742,0.00058019435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074260174,0.0028911203,0.0018258048,0.009011217,0.0011384257,0.002953183,0.0034825935,0.0022854828,0.0027474312],"category_scores_gemma":[0.018389748,0.0008650042,0.0023631083,0.0064134644,0.0006440734,0.004123918,0.0015310817,0.0030099386,0.0043552634],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012196966,0.00090305676,0.022558436,0.003831805,0.0012188558,0.0004990508,0.0005540116,0.11670359,0.0045139324,0.007347565,0.116210304,0.7244398],"study_design_scores_gemma":[0.000088032575,0.0004986486,0.006279434,0.0006601429,0.00032636125,0.00048296954,0.00041873293,0.91324735,0.0037935113,0.013146258,0.06093455,0.00012404169],"about_ca_topic_score_codex":0.02006884,"about_ca_topic_score_gemma":0.02870511,"teacher_disagreement_score":0.02006884,"about_ca_system_score_codex":0.002439863,"about_ca_system_score_gemma":0.0035893961,"threshold_uncertainty_score":0.039904},"labels":[],"label_agreement":null},{"id":"W4220799018","doi":"10.1016/j.artmed.2022.102282","title":"Chinese clinical named entity recognition via multi-head self-attention based BiLSTM-CRF","year":2022,"lang":"en","type":"article","venue":"Artificial Intelligence in Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":132,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Science Foundation of Hunan Province","keywords":"Computer science; Conditional random field; Artificial intelligence; Named-entity recognition; Natural language processing; Benchmark (surveying); Task (project management); Feature (linguistics); Deep learning; Embedding; Character (mathematics); Pattern recognition (psychology)","score_opus":0.14290939386940332,"score_gpt":0.40900090944707085,"score_spread":0.2660915155776675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220799018","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15827504,0.008927933,0.77295417,0.0028218785,0.0016115783,0.0005935335,0.027811948,0.01741393,0.009589937],"genre_scores_gemma":[0.7642004,0.0019109317,0.18242024,0.000746178,0.0007933076,0.0003955432,0.040352058,0.00048725598,0.008694112],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986059,0.00024677103,0.00017800849,0.0006542933,0.00014345488,0.00017160428],"domain_scores_gemma":[0.9981864,0.0007901746,0.0001529721,0.00025692311,0.00052588305,0.00008758158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018586586,0.0010445656,0.0011082423,0.002349488,0.00079991773,0.0007732113,0.0013222214,0.0013847663,0.0041276305],"category_scores_gemma":[0.0032564618,0.00044844853,0.0015301746,0.0025079553,0.0003424111,0.0017002797,0.0013621425,0.0014972199,0.0030811294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013372581,0.00050811516,0.034867436,0.0012609741,0.0005031955,0.0016421003,0.000500135,0.03502751,0.030360743,0.005582504,0.09148392,0.79692614],"study_design_scores_gemma":[0.000115633644,0.0002467635,0.029870132,0.00012197145,0.00084209477,0.0017131016,0.00017624663,0.90478617,0.024343383,0.0120411,0.025591629,0.00015174063],"about_ca_topic_score_codex":0.014314185,"about_ca_topic_score_gemma":0.018229946,"teacher_disagreement_score":0.014314185,"about_ca_system_score_codex":0.0007731841,"about_ca_system_score_gemma":0.0024038272,"threshold_uncertainty_score":0.028461695},"labels":[],"label_agreement":null},{"id":"W4220896903","doi":"10.1145/3507356","title":"Leveraging Narrative to Generate Movie Script","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada); Université de Montréal","funders":"","keywords":"Computer science; Narrative; Scripting language; Context (archaeology); Upload; Unavailability; Task (project management); Construct (python library); Information retrieval; Artificial intelligence; Natural language processing; World Wide Web; Linguistics; Programming language","score_opus":0.03357747367188199,"score_gpt":0.24551363434961782,"score_spread":0.21193616067773582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220896903","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19913226,0.0055119353,0.6519421,0.0023979913,0.0007547352,0.0031516596,0.06144606,0.058506146,0.01715713],"genre_scores_gemma":[0.28473267,0.0010243439,0.59312737,0.0006304219,0.00015959215,0.0008682896,0.10859724,0.0013781668,0.009481948],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99888605,0.000376968,0.00011516122,0.00041126923,0.0001586242,0.000051921194],"domain_scores_gemma":[0.99771047,0.0011632855,0.00021599985,0.00045820716,0.0003121934,0.00013987662],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001006298,0.0015474637,0.0006471928,0.0014175086,0.00046062324,0.001031738,0.0012862239,0.0011058089,0.003305051],"category_scores_gemma":[0.007260557,0.0003497479,0.0009467202,0.00083477644,0.00030057394,0.0023963342,0.00081151095,0.0010070878,0.0040784692],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013509454,0.00084435346,0.021741437,0.003527634,0.00026010856,0.002078766,0.0020966479,0.040653985,0.060608394,0.008855675,0.13302061,0.7249614],"study_design_scores_gemma":[0.0002559425,0.0008183273,0.0111497445,0.00028747888,0.00016586274,0.0018874931,0.0014343271,0.71079725,0.0683767,0.016668884,0.18797837,0.00017963773],"about_ca_topic_score_codex":0.003101846,"about_ca_topic_score_gemma":0.0073796026,"teacher_disagreement_score":0.003305051,"about_ca_system_score_codex":0.0006596388,"about_ca_system_score_gemma":0.0007117531,"threshold_uncertainty_score":0.011056483},"labels":[],"label_agreement":null},{"id":"W4220937440","doi":"10.1002/essoar.10510978.1","title":"m-NLP inference models using simulation and regression techniques","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Norges Forskningsråd; H2020 European Research Council; China Scholarship Council; Compute Canada","keywords":"Preprint; Inference; Computer science; World Wide Web; Information retrieval; Natural language processing; Artificial intelligence","score_opus":0.1589566516328696,"score_gpt":0.38556881623716716,"score_spread":0.22661216460429756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220937440","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008896817,0.00080402807,0.9831549,0.0013170129,0.0001456238,0.00013786489,0.00136701,0.0018425988,0.0023342152],"genre_scores_gemma":[0.32198396,0.0016047147,0.650956,0.0008412293,0.00066057,0.0015374657,0.008561212,0.0015409733,0.012313898],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912032,0.006437956,0.00030260455,0.0011697813,0.000592345,0.00029413417],"domain_scores_gemma":[0.92731744,0.065145455,0.0015100613,0.0030876503,0.0024415986,0.0004978564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013539758,0.0015768732,0.0023102094,0.002344274,0.0013040433,0.0032772976,0.0037656534,0.0027261828,0.010606596],"category_scores_gemma":[0.0727993,0.0016061517,0.0028911242,0.0030182265,0.0014704269,0.0040344144,0.0025122832,0.005145498,0.0037419382],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031405804,0.00015035698,0.0048155407,0.00033126757,0.0004675619,0.00021604337,0.00022245599,0.7917653,0.00027841044,0.0967257,0.016187085,0.08852624],"study_design_scores_gemma":[0.0000328843,0.000010046837,0.00016540871,0.000024021987,0.000021470483,0.000013115409,0.00001512895,0.95517397,0.00011566987,0.04302121,0.00139685,0.000010234856],"about_ca_topic_score_codex":0.029214384,"about_ca_topic_score_gemma":0.02790926,"teacher_disagreement_score":0.029214384,"about_ca_system_score_codex":0.0020911482,"about_ca_system_score_gemma":0.0036714168,"threshold_uncertainty_score":0.07160598},"labels":[],"label_agreement":null},{"id":"W4221016102","doi":"10.1007/s00521-022-07072-0","title":"Mutually improved dense retriever and GNN-based reader for arbitrary-hop open-domain question answering","year":2022,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Labrador Retriever; Question answering; Open domain; Asynchronous communication; Leverage (statistics); Graph; Information retrieval; Preprocessor; Artificial intelligence; Theoretical computer science; Computer network","score_opus":0.02400630599384165,"score_gpt":0.28512810381511394,"score_spread":0.2611217978212723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221016102","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03147912,0.0021931978,0.9281663,0.0006970187,0.0004965766,0.00028764253,0.002094486,0.025457518,0.00912806],"genre_scores_gemma":[0.34573185,0.0007952334,0.61517024,0.0011866589,0.0006260509,0.0003155052,0.011283897,0.0016010305,0.023289507],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977016,0.0004644805,0.00016137895,0.0007397859,0.0006469505,0.00028581094],"domain_scores_gemma":[0.996869,0.0010009112,0.000101835736,0.0012917215,0.00059431064,0.0001422444],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015581086,0.001519378,0.0026672678,0.0021895866,0.00097124523,0.0021523836,0.003361476,0.002649359,0.014573085],"category_scores_gemma":[0.007186946,0.0006275797,0.0014954873,0.0018145846,0.0009799453,0.006050974,0.0045449273,0.0021633203,0.011457606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011358768,0.0006369547,0.001619713,0.0006909391,0.00023260392,0.0004988037,0.00043699416,0.027310701,0.030017953,0.028326975,0.052080303,0.85701215],"study_design_scores_gemma":[0.000200665,0.00027031178,0.0008687958,0.000058625013,0.00022884995,0.0008061278,0.00029897472,0.87833107,0.024293276,0.07383833,0.020700727,0.00010423874],"about_ca_topic_score_codex":0.004713333,"about_ca_topic_score_gemma":0.011231373,"teacher_disagreement_score":0.014573085,"about_ca_system_score_codex":0.0007594115,"about_ca_system_score_gemma":0.002022084,"threshold_uncertainty_score":0.04875189},"labels":[],"label_agreement":null},{"id":"W4221039763","doi":"10.18438/eblip30014","title":"Natural Language Processing for Virtual Reference Analysis","year":2022,"lang":"en","type":"article","venue":"Evidence Based Library and Information Practice","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"University of Toronto Scarborough; University of Toronto","keywords":"Computer science; Toolchain; Information retrieval; Python (programming language); Coding (social sciences); Natural language processing; World Wide Web; Artificial intelligence; Programming language; Software","score_opus":0.017753414837502667,"score_gpt":0.27309779600051143,"score_spread":0.2553443811630088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221039763","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025616493,0.00052478106,0.96084267,0.0010809726,0.00030388337,0.0009023986,0.007721691,0.016403852,0.009658039],"genre_scores_gemma":[0.029100828,0.00041657418,0.94487,0.0004197842,0.00016642932,0.002456896,0.016047705,0.0024102742,0.0041115903],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.966896,0.016141241,0.0036762387,0.005518004,0.0072035384,0.0005649802],"domain_scores_gemma":[0.9246473,0.039013296,0.0047512855,0.015686862,0.015218851,0.0006824022],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.022640267,0.002019549,0.0012331376,0.009419349,0.002425385,0.007770438,0.0033786155,0.00143378,0.039006054],"category_scores_gemma":[0.094125256,0.0009406055,0.0030125938,0.007862377,0.0027442337,0.007158913,0.0068158777,0.003057158,0.024780327],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003314145,0.0001603319,0.002877127,0.0037437037,0.00024058754,0.0008932095,0.0058271633,0.0076732677,0.015780555,0.15496941,0.09662605,0.71087724],"study_design_scores_gemma":[0.00008485593,0.00015803329,0.00405293,0.0012860883,0.000101814585,0.00077942584,0.0029714275,0.07581764,0.018258106,0.34833443,0.54787534,0.00027989148],"about_ca_topic_score_codex":0.005981767,"about_ca_topic_score_gemma":0.005318165,"teacher_disagreement_score":0.9922296,"about_ca_system_score_codex":0.0032418498,"about_ca_system_score_gemma":0.0067870608,"threshold_uncertainty_score":0.13048828},"labels":[],"label_agreement":null},{"id":"W4221053465","doi":"10.1162/tacl_a_00458","title":"Neuro-symbolic Natural Logic with Introspective Revision for Natural Language Inference","year":2022,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Interpretability; Computer science; Artificial intelligence; Inference; Generalization; Introspection; Spurious relationship; Machine learning; Rule of inference; Natural language; Natural (archaeology); Overfitting; Artificial neural network; Cognitive psychology","score_opus":0.01236504266629197,"score_gpt":0.2721111208052126,"score_spread":0.25974607813892064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221053465","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020584852,0.00027885425,0.9748824,0.0004582441,0.000036207643,0.000061128696,0.00014010858,0.0012879563,0.0022701996],"genre_scores_gemma":[0.7641208,0.00022189441,0.23333044,0.00020865757,0.00005914608,0.00015733199,0.00026786106,0.00011301127,0.0015207307],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99874544,0.00051897427,0.00007245454,0.0002496143,0.00032022604,0.00009310909],"domain_scores_gemma":[0.99690944,0.0017125926,0.00028936792,0.0006043874,0.00037113644,0.00011309936],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018431821,0.0005528256,0.00063372304,0.0008508851,0.00045766545,0.0014039811,0.0017927582,0.0006745504,0.0031331866],"category_scores_gemma":[0.0066421763,0.00033350286,0.0011320383,0.0005914749,0.0017912085,0.0027184607,0.0015744091,0.0020966379,0.00040301945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020577153,0.00018852933,0.0017939106,0.00020107446,0.00014037311,0.00031224728,0.0003263401,0.6272089,0.0048861885,0.20202866,0.0027422162,0.15996574],"study_design_scores_gemma":[0.00001037496,0.000015300271,0.00009419651,0.000008019105,0.000009638651,0.0000228229,0.0000070465367,0.9390288,0.0006276523,0.059622042,0.00054721313,0.0000069128746],"about_ca_topic_score_codex":0.005793418,"about_ca_topic_score_gemma":0.008246122,"teacher_disagreement_score":0.005793418,"about_ca_system_score_codex":0.0016618416,"about_ca_system_score_gemma":0.0020159255,"threshold_uncertainty_score":0.012057543},"labels":[],"label_agreement":null},{"id":"W4221147223","doi":"10.18653/v1/2022.acl-long.359","title":"E-LANG: Energy-Based Joint Inferencing of Super and Swift Language Models","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Swift; Inference; Language model; Computation; Code (set theory); Encoder; Set (abstract data type); Artificial intelligence; Sequence (biology); Architecture; Energy (signal processing); Computer engineering; Machine learning; Theoretical computer science; Programming language; Operating system","score_opus":0.011359810896829392,"score_gpt":0.21521957003745917,"score_spread":0.2038597591406298,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221147223","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016142182,0.00020369972,0.96336645,0.00026964015,0.00008907574,0.0000606321,0.00049391476,0.016692849,0.002681524],"genre_scores_gemma":[0.3616567,0.00024898554,0.6214909,0.00054790557,0.00012079851,0.00019182776,0.0039688656,0.0033011958,0.008472863],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990081,0.00027511772,0.000048563987,0.00026851098,0.00028101951,0.000118603675],"domain_scores_gemma":[0.9982681,0.0007220227,0.000092385635,0.0005931875,0.00023682366,0.00008746171],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016594866,0.0015310841,0.0009943703,0.00079235225,0.00056606665,0.0016584015,0.0026711763,0.0011592115,0.006691596],"category_scores_gemma":[0.0058573587,0.0011176255,0.0012221182,0.00060641434,0.0008327049,0.0048056347,0.0028843116,0.0030436474,0.0033995206],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007081508,0.00033995404,0.0027626897,0.00024570065,0.00026183538,0.0004580556,0.00036316045,0.380598,0.020320771,0.042873673,0.02387752,0.52719045],"study_design_scores_gemma":[0.000024216772,0.000028729786,0.00010810037,0.00000748452,0.000015027691,0.000035562352,0.000023228999,0.9738341,0.0050149765,0.018724231,0.0021700927,0.000014219169],"about_ca_topic_score_codex":0.005619434,"about_ca_topic_score_gemma":0.016091969,"teacher_disagreement_score":0.006691596,"about_ca_system_score_codex":0.0008416774,"about_ca_system_score_gemma":0.0016454565,"threshold_uncertainty_score":0.022385657},"labels":[],"label_agreement":null},{"id":"W4221148028","doi":"10.1007/s11633-022-1387-3","title":"EVA2.0: Investigating Open-domain Chinese Dialogue Systems with Large-scale Pre-training","year":2023,"lang":"en","type":"article","venue":"Machine Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Open domain; Domain (mathematical analysis); Chatbot; Scale (ratio); Quality (philosophy); Key (lock); Open research; Architecture; Artificial intelligence; Code (set theory); Data science; World Wide Web; Question answering; Computer security","score_opus":0.13273794375393888,"score_gpt":0.41194765496336117,"score_spread":0.2792097112094223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221148028","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97663385,0.0002402836,0.014925567,0.00019273828,0.00007105594,0.00019757179,0.0016754674,0.0017581636,0.004305343],"genre_scores_gemma":[0.9780618,0.00006271251,0.014278415,0.000052685635,0.000027898528,0.00021121759,0.004617118,0.00012932817,0.002558779],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99885595,0.0006460803,0.000045507462,0.0002268074,0.000118422184,0.000107243366],"domain_scores_gemma":[0.99401337,0.004580474,0.00013123677,0.0005182947,0.00044490074,0.00031157097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024787537,0.0007537582,0.00055468315,0.0008701485,0.0011088818,0.0011378629,0.0012812755,0.000770549,0.0030871183],"category_scores_gemma":[0.0076097036,0.00031018557,0.00040779388,0.0008066284,0.0006403853,0.0014511027,0.0014054058,0.0013606715,0.000949384],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0058391443,0.005844142,0.11169517,0.0027846564,0.00086671003,0.0026610747,0.02314947,0.20296052,0.12622209,0.014345845,0.048382405,0.4552489],"study_design_scores_gemma":[0.00032229806,0.0012515322,0.073920354,0.000054219512,0.0001979192,0.000372756,0.0053888005,0.8715513,0.030071797,0.004830266,0.0118703,0.00016843608],"about_ca_topic_score_codex":0.019575171,"about_ca_topic_score_gemma":0.020996366,"teacher_disagreement_score":0.019575171,"about_ca_system_score_codex":0.00084250903,"about_ca_system_score_gemma":0.0011684763,"threshold_uncertainty_score":0.03892249},"labels":[],"label_agreement":null},{"id":"W4223491992","doi":"10.18653/v1/2022.bionlp-1.2","title":"A sequence-to-sequence approach for document-level relation extraction","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Princess Margaret Cancer Centre; Lunenfeld-Tanenbaum Research Institute; University Health Network; Vector Institute; University of Toronto","funders":"National Institutes of Health; Compute Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Pipeline (software); Coreference; Relationship extraction; Sequence (biology); Task (project management); Sentence; Relation (database); Information extraction; Natural language processing; Code (set theory); Information retrieval; Artificial intelligence; Named-entity recognition; Data mining; Resolution (logic); Programming language; Set (abstract data type)","score_opus":0.13796717483022905,"score_gpt":0.32693444500118835,"score_spread":0.1889672701709593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4223491992","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003757854,0.00085222913,0.96033555,0.0003313631,0.00020614525,0.00042052145,0.0060309116,0.02576833,0.0022972045],"genre_scores_gemma":[0.033376265,0.00061824307,0.9299126,0.00031302078,0.00019384237,0.00049153087,0.02640117,0.0011125345,0.007580804],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979995,0.0003540422,0.00017678787,0.00097328564,0.00038321014,0.00011320575],"domain_scores_gemma":[0.9967577,0.0012029642,0.00023889817,0.00082367496,0.00083719747,0.00013952679],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020909929,0.0030513143,0.0012857462,0.0054157358,0.0011551126,0.0015868935,0.0021945273,0.0020233728,0.010887267],"category_scores_gemma":[0.0053488216,0.0011051601,0.002575161,0.005007682,0.0006549359,0.004502529,0.002709861,0.0031558461,0.015203912],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039848065,0.00043914004,0.0024229977,0.0009043445,0.00027757717,0.0005244885,0.0005741406,0.017331375,0.048889417,0.014937668,0.087408274,0.82589215],"study_design_scores_gemma":[0.00011299822,0.0004619924,0.0036395986,0.0001761823,0.0002916074,0.0017125024,0.00037892227,0.6750948,0.058683712,0.07637356,0.18288025,0.00019392106],"about_ca_topic_score_codex":0.007126076,"about_ca_topic_score_gemma":0.018193206,"teacher_disagreement_score":0.010887267,"about_ca_system_score_codex":0.0010537178,"about_ca_system_score_gemma":0.0029704592,"threshold_uncertainty_score":0.036421537},"labels":[],"label_agreement":null},{"id":"W4224226311","doi":"10.2196/35606","title":"Multi-Label Classification in Patient-Doctor Dialogues With the RoBERTa-WWM-ext + CNN (Robustly Optimized Bidirectional Encoder Representations From Transformers Pretraining Approach With Whole Word Masking Extended Combining a Convolutional Neural Network) Model: Named Entity Study","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Key Research and Development Program of China","keywords":"Computer science; Encoder; Sentence; Transformer; Natural language processing; Artificial intelligence; Convolutional neural network; Chatbot; F1 score","score_opus":0.060152736999827976,"score_gpt":0.2867855212292952,"score_spread":0.2266327842294672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224226311","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5054017,0.00079303974,0.48362103,0.00087236083,0.00022957227,0.00020433111,0.0007490905,0.004190204,0.003938764],"genre_scores_gemma":[0.92619044,0.00009716745,0.06716695,0.00018854626,0.000072376315,0.000089483096,0.00095747266,0.000059572107,0.005177974],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995503,0.00009918699,0.000024705821,0.00018745854,0.00005734063,0.00008110614],"domain_scores_gemma":[0.9994831,0.00023204922,0.000061450286,0.000059943104,0.00011195439,0.000051414045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000932281,0.0007716689,0.00037938237,0.00058751384,0.00037613197,0.00051713246,0.000957719,0.0010762024,0.00145595],"category_scores_gemma":[0.0014357994,0.0002761521,0.0007762501,0.00033298563,0.00029073074,0.00092966843,0.00076795806,0.001050648,0.00063226616],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014951894,0.0009737265,0.029190844,0.00021408992,0.00023147889,0.0009934979,0.0010338322,0.18866387,0.05057481,0.003211865,0.008598766,0.71481806],"study_design_scores_gemma":[0.0000076929355,0.00007659295,0.001959355,0.000004910021,0.00002724414,0.000041448006,0.000048679496,0.9891922,0.007370735,0.00068511814,0.0005744296,0.00001159696],"about_ca_topic_score_codex":0.007984833,"about_ca_topic_score_gemma":0.011380154,"teacher_disagreement_score":0.007984833,"about_ca_system_score_codex":0.00087726937,"about_ca_system_score_gemma":0.0007208315,"threshold_uncertainty_score":0.01587671},"labels":[],"label_agreement":null},{"id":"W4224278993","doi":"10.3390/info13040205","title":"Medical Knowledge Graph Completion Based on Word Embeddings","year":2022,"lang":"en","type":"article","venue":"Information","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"Natural Science Foundation of Beijing Municipality","keywords":"Word2vec; Computer science; RDF; Knowledge graph; Terminology; Information retrieval; Word (group theory); Natural language processing; Semantics (computer science); Graph; Artificial intelligence; Theoretical computer science; Semantic Web; Embedding; Mathematics","score_opus":0.0161672622546258,"score_gpt":0.2530601875034533,"score_spread":0.23689292524882752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224278993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.083848506,0.0017731644,0.8957277,0.0011744751,0.00020646014,0.00029036548,0.0077811535,0.005925679,0.0032725052],"genre_scores_gemma":[0.5900214,0.0014445463,0.36093232,0.0006910027,0.0001948455,0.0004929561,0.040386416,0.0005295525,0.0053069494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99912935,0.000139239,0.000078438105,0.00043990707,0.00014425765,0.00006875505],"domain_scores_gemma":[0.9986877,0.0005788027,0.00017734525,0.00022504354,0.00026825824,0.000062854895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005763983,0.0012658674,0.00065380963,0.0029102126,0.00038477228,0.0008352753,0.00092958286,0.00093863416,0.002755174],"category_scores_gemma":[0.0043344605,0.00035530614,0.001216582,0.0020834496,0.00070182054,0.0034004124,0.0012585147,0.001478209,0.0017233856],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005350942,0.00029755503,0.015696276,0.0009023241,0.00031398778,0.0009029443,0.00088401843,0.1162316,0.0156334,0.026066262,0.030818539,0.79171795],"study_design_scores_gemma":[0.00005248659,0.00017671553,0.0035270462,0.00013442463,0.00012937565,0.00062517164,0.00041050953,0.90395826,0.009856975,0.062196262,0.018874485,0.00005822138],"about_ca_topic_score_codex":0.0076789437,"about_ca_topic_score_gemma":0.013016808,"teacher_disagreement_score":0.0076789437,"about_ca_system_score_codex":0.00071052444,"about_ca_system_score_gemma":0.0011753606,"threshold_uncertainty_score":0.015268505},"labels":[],"label_agreement":null},{"id":"W4224279631","doi":"10.18653/v1/2022.naacl-main.426","title":"Less is More: Learning to Refine Dialogue History for Personalized Dialogue Generation","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Renmin University of China; National Natural Science Foundation of China","keywords":"Computer science; Computational linguistics; Human language; Natural language processing; Association (psychology); Artificial intelligence; Cognitive science; Linguistics; Psychology; Philosophy; Epistemology","score_opus":0.03673297298027899,"score_gpt":0.2583922274011966,"score_spread":0.2216592544209176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224279631","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.063619286,0.0026414741,0.89985496,0.0012230539,0.0005296387,0.00052558223,0.0014358269,0.022569472,0.0076007806],"genre_scores_gemma":[0.6413783,0.0005269369,0.34505394,0.0005945992,0.0002594467,0.00063302205,0.0037113999,0.0012170324,0.0066253995],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979151,0.0008988633,0.000086241176,0.00074978464,0.00020020606,0.00014982848],"domain_scores_gemma":[0.99364036,0.004583235,0.00022191089,0.0006880116,0.00053608767,0.0003304275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026864237,0.001797343,0.0010447179,0.0014845344,0.0009313307,0.0014179089,0.0022604514,0.0015481912,0.009601868],"category_scores_gemma":[0.014704518,0.0009439407,0.00081110233,0.00069449114,0.0007053992,0.0043028328,0.0028319992,0.0019043963,0.003241534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014286208,0.0004309918,0.0037973193,0.00040043602,0.00011669948,0.00015602224,0.0014485907,0.021197472,0.010511891,0.0037503757,0.022944026,0.9338176],"study_design_scores_gemma":[0.00037559698,0.0004250806,0.0026846686,0.00010876019,0.00019893357,0.00020551987,0.0006882646,0.93938136,0.011621889,0.03158318,0.012641003,0.00008579755],"about_ca_topic_score_codex":0.0038643675,"about_ca_topic_score_gemma":0.007213588,"teacher_disagreement_score":0.009601868,"about_ca_system_score_codex":0.00077438384,"about_ca_system_score_gemma":0.001077656,"threshold_uncertainty_score":0.03212142},"labels":[],"label_agreement":null},{"id":"W4224325974","doi":"10.1093/database/baac069","title":"Multi-label classification for biomedical literature: an overview of the BioCreative VII LitCovid Track for COVID-19 literature topic annotations","year":2022,"lang":"en","type":"article","venue":"Database","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"U.S. National Library of Medicine; Medical Research Council; National Institutes of Health","keywords":"Computer science; Coronavirus disease 2019 (COVID-19); Annotation; Named-entity recognition; Scientific literature; Information retrieval; Data science; Artificial intelligence; Natural language processing; World Wide Web; Task (project management); Medicine; Infectious disease (medical specialty)","score_opus":0.2046024088696558,"score_gpt":0.4089928303778776,"score_spread":0.2043904215082218,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224325974","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012055004,0.073803544,0.1060613,0.011625832,0.0045686616,0.003988834,0.6112822,0.14208667,0.034527957],"genre_scores_gemma":[0.0055425055,0.0098674055,0.15123805,0.0030053922,0.00079790497,0.0022309653,0.8128781,0.006275997,0.008163751],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9812501,0.003814021,0.0032635592,0.004289219,0.00638824,0.0009948596],"domain_scores_gemma":[0.955544,0.015401774,0.0046804706,0.007638965,0.012559704,0.0041751866],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.019823898,0.0028758228,0.0031121548,0.042533647,0.004193842,0.011486439,0.0051621133,0.0036162399,0.022048146],"category_scores_gemma":[0.047683895,0.0016000797,0.004641766,0.03224641,0.0011904661,0.008785834,0.009696248,0.003386465,0.039940547],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054124574,0.00032313002,0.0058385152,0.010145025,0.00042799243,0.00039076182,0.0007372942,0.000972156,0.010533012,0.0032573398,0.73061574,0.2362178],"study_design_scores_gemma":[0.00014299745,0.00017197097,0.007483691,0.0021608411,0.00016065376,0.00049917685,0.00025668746,0.0045755873,0.004261619,0.0032111774,0.9769165,0.00015915488],"about_ca_topic_score_codex":0.012929982,"about_ca_topic_score_gemma":0.031409204,"teacher_disagreement_score":0.95746636,"about_ca_system_score_codex":0.0047586537,"about_ca_system_score_gemma":0.01145212,"threshold_uncertainty_score":0.10484004},"labels":[],"label_agreement":null},{"id":"W4225102708","doi":"10.18653/v1/2022.naacl-main.212","title":"Document-Level Relation Extraction with Sentences Importance Estimation and Focusing","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Key Research and Development Program of China; DeepMind; Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Sentence; Computer science; Relationship extraction; Focus (optics); Natural language processing; Artificial intelligence; Relation (database); Variety (cybernetics); Graph; Sequence (biology); Information extraction; Data mining; Theoretical computer science","score_opus":0.019314828825849745,"score_gpt":0.2540263137987291,"score_spread":0.23471148497287933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225102708","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03179134,0.003124836,0.93772537,0.0007372559,0.0002234489,0.00040339772,0.005060892,0.017548168,0.003385365],"genre_scores_gemma":[0.25839922,0.0014959816,0.7058355,0.00050318206,0.000552157,0.00036836736,0.02396233,0.00069315603,0.008190114],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984029,0.00027006818,0.00014687708,0.0006351795,0.00041593643,0.00012906153],"domain_scores_gemma":[0.99715304,0.0012615494,0.0003176764,0.00057790615,0.0005899732,0.00009993167],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018578079,0.0026444674,0.001311423,0.0076797907,0.0006969177,0.0013728896,0.0019194868,0.0018711977,0.0028541514],"category_scores_gemma":[0.00498762,0.000555238,0.0018777907,0.0037764895,0.00051674567,0.0039671436,0.0016715188,0.002049094,0.004873831],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041757713,0.00030793337,0.0076444233,0.00047627836,0.00021920267,0.00075165933,0.00031084675,0.014647095,0.043812063,0.0050732386,0.039935302,0.88640434],"study_design_scores_gemma":[0.00009624665,0.00069609366,0.014494562,0.00015322826,0.0005084189,0.0023915104,0.00033223233,0.82550883,0.085620336,0.02576863,0.04424968,0.0001801461],"about_ca_topic_score_codex":0.0031181744,"about_ca_topic_score_gemma":0.006146336,"teacher_disagreement_score":0.0076797907,"about_ca_system_score_codex":0.0008169941,"about_ca_system_score_gemma":0.0009870906,"threshold_uncertainty_score":0.00982511},"labels":[],"label_agreement":null},{"id":"W4225334559","doi":"10.21105/joss.04038","title":"ADaPT-ML: A Data Programming Template for MachineLearning","year":2022,"lang":"en","type":"article","venue":"The Journal of Open Source Software","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Mitacs","keywords":"Computer science; Task (project management); Class (philosophy); Machine learning; Domain (mathematical analysis); Point (geometry); Labeled data; Artificial intelligence; Training set; Data type; Software deployment; Software engineering; Programming language","score_opus":0.09902702177909116,"score_gpt":0.3346960179557344,"score_spread":0.23566899617664322,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225334559","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002784109,0.000055658023,0.9742556,0.00020155935,0.00008330515,0.000075084434,0.0009854676,0.022574445,0.0014905134],"genre_scores_gemma":[0.014718635,0.00016071514,0.9660366,0.0005260731,0.00013648422,0.00082914205,0.0034056231,0.009700862,0.0044858507],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971506,0.0008026091,0.0003340993,0.00072770094,0.0008286088,0.00015631643],"domain_scores_gemma":[0.9945016,0.003072341,0.00019828521,0.0014254523,0.0006458523,0.00015643107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041665826,0.0017492182,0.001175471,0.0013485594,0.00071226264,0.003676278,0.0051067495,0.0022852055,0.02443701],"category_scores_gemma":[0.019055475,0.001302034,0.0030206365,0.0017842903,0.0011827981,0.0042852038,0.0045977635,0.006415279,0.02099271],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005944629,0.00028329046,0.001916226,0.00070768606,0.00026500807,0.00035379207,0.00042331454,0.07422519,0.0051350044,0.13331561,0.15594873,0.6268317],"study_design_scores_gemma":[0.0001331171,0.00006361375,0.00027397467,0.000108244065,0.000045622706,0.00026702078,0.000057266927,0.60840464,0.012163081,0.23863122,0.13978943,0.00006269635],"about_ca_topic_score_codex":0.0015917275,"about_ca_topic_score_gemma":0.003336636,"teacher_disagreement_score":0.02443701,"about_ca_system_score_codex":0.00086167903,"about_ca_system_score_gemma":0.0015487117,"threshold_uncertainty_score":0.081749916},"labels":[],"label_agreement":null},{"id":"W4225362522","doi":"10.1007/978-3-030-99739-7_24","title":"How Different are Pre-trained Transformers for Text Ranking?","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Canadian Institute of Steel Construction","keywords":"Computer science; Transformer; Artificial intelligence; Ranking (information retrieval); Machine learning; Information retrieval; Relevance (law); Question answering; Encoder; Task (project management); Recall; Artificial neural network; Learning to rank; Precision and recall; Deep learning; Natural language processing; Deep neural networks","score_opus":0.02048011377708883,"score_gpt":0.23805279336187277,"score_spread":0.21757267958478393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225362522","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10479865,0.009184176,0.8302152,0.006971478,0.0025278316,0.00038140136,0.003716363,0.025811512,0.016393345],"genre_scores_gemma":[0.7054167,0.0030310408,0.26223588,0.0019765291,0.001187335,0.00029951253,0.009303985,0.00401352,0.012535555],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99252397,0.0033312507,0.00050554785,0.0017414502,0.001191487,0.00070626667],"domain_scores_gemma":[0.985658,0.0076769954,0.00033157913,0.0034135715,0.002274496,0.00064535474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008871355,0.0017449288,0.0017714191,0.002509979,0.00091912405,0.0046997773,0.0026406392,0.0029160744,0.007967296],"category_scores_gemma":[0.04243112,0.0008516066,0.0015399093,0.0019110771,0.0011318373,0.016683428,0.0016900572,0.0043965233,0.011557779],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018336232,0.00039080935,0.0075336797,0.00047208203,0.0005076202,0.00005599829,0.00018252817,0.01421436,0.011523782,0.012201909,0.04174099,0.9093427],"study_design_scores_gemma":[0.00076445384,0.0013222303,0.0109126195,0.0004389734,0.0009266257,0.00065133616,0.0010639795,0.7102238,0.06508983,0.1709856,0.037322722,0.00029783876],"about_ca_topic_score_codex":0.004799989,"about_ca_topic_score_gemma":0.009493167,"teacher_disagreement_score":0.008871355,"about_ca_system_score_codex":0.0013749878,"about_ca_system_score_gemma":0.0026689097,"threshold_uncertainty_score":0.046916783},"labels":[],"label_agreement":null},{"id":"W4225552142","doi":"10.1145/3527546.3527570","title":"Report on the 15th round of NII testbeds and community for information access research (NTCIR-15)","year":2021,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Task (project management); Scope (computer science); Question answering; Information retrieval; World Wide Web; Artificial intelligence; Programming language","score_opus":0.1691181348945088,"score_gpt":0.39618905820630096,"score_spread":0.22707092331179216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225552142","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022969997,0.01111493,0.047278438,0.09264924,0.078964785,0.02329331,0.46967584,0.020082952,0.23397048],"genre_scores_gemma":[0.040669706,0.0027369265,0.062025405,0.016880829,0.004665878,0.019339537,0.60907197,0.0051172175,0.23949255],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9805787,0.0056169904,0.0007789745,0.0016119581,0.007575648,0.003837603],"domain_scores_gemma":[0.91539097,0.0058100997,0.0017532003,0.0057064667,0.033410016,0.03792934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.048528258,0.0023343728,0.0021943757,0.004811707,0.005313405,0.009853782,0.004215344,0.004461051,0.086294904],"category_scores_gemma":[0.02938752,0.0011141209,0.0020521346,0.0034282028,0.0012791838,0.006166449,0.013661634,0.0051575336,0.089781664],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032692053,0.00020306434,0.0009256289,0.00019870172,0.00002517222,0.000068510126,0.0001534151,0.00015373468,0.00093477854,0.0015344935,0.98213965,0.013336018],"study_design_scores_gemma":[0.00028420167,0.00038679966,0.007978342,0.00019305569,0.00005036005,0.000059511378,0.00070034625,0.0006589511,0.0018911161,0.0017658664,0.9859476,0.00008382347],"about_ca_topic_score_codex":0.04794457,"about_ca_topic_score_gemma":0.067600675,"teacher_disagreement_score":0.086294904,"about_ca_system_score_codex":0.0051865233,"about_ca_system_score_gemma":0.025015388,"threshold_uncertainty_score":0.2886852},"labels":[],"label_agreement":null},{"id":"W4225580830","doi":"10.18653/v1/2022.acl-long.396","title":"Subgraph Retrieval Enhanced Model for Multi-hop Knowledge Base Question Answering","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute","funders":"National Natural Science Foundation of China; Tencent","keywords":"Computer science; Semantic reasoner; Information retrieval; Embedding; Question answering; Knowledge base; Heuristic; Process (computing); Pruning; Base (topology); Artificial intelligence; Programming language; Mathematics","score_opus":0.01924940758376219,"score_gpt":0.2685602883043876,"score_spread":0.2493108807206254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225580830","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011350178,0.0008179715,0.9792336,0.00043196947,0.000057145095,0.00013915072,0.0011880073,0.0050016902,0.0017802147],"genre_scores_gemma":[0.4220194,0.00093943154,0.5544519,0.0007347412,0.00017112226,0.0006431076,0.009745024,0.00059601164,0.010699269],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993967,0.00015244326,0.000034355402,0.00024788195,0.000116297,0.000052356063],"domain_scores_gemma":[0.99923646,0.00039315308,0.00004310509,0.00014231433,0.00014744507,0.000037514623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008900126,0.0009895765,0.00092107855,0.0014074665,0.0004246868,0.00084098714,0.0025093239,0.0018353158,0.0055850344],"category_scores_gemma":[0.0031885644,0.0004097981,0.0015261386,0.0011152831,0.000564307,0.0022981479,0.0014078483,0.0015886917,0.0026714723],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004056395,0.00028842475,0.0018916163,0.0005414544,0.00020884057,0.00036179263,0.00051118026,0.5027993,0.016302973,0.037365846,0.024983179,0.4143398],"study_design_scores_gemma":[0.0000147035435,0.000028832845,0.00015475179,0.000008717163,0.0000269396,0.000041607527,0.00001543558,0.9807989,0.0010719922,0.015620594,0.0022094876,0.000008006759],"about_ca_topic_score_codex":0.010415064,"about_ca_topic_score_gemma":0.015400168,"teacher_disagreement_score":0.010415064,"about_ca_system_score_codex":0.0011149686,"about_ca_system_score_gemma":0.0010366194,"threshold_uncertainty_score":0.020708919},"labels":[],"label_agreement":null},{"id":"W4225656840","doi":"10.1007/978-3-030-99739-7_21","title":"GameOfThronesQA: Answer-Aware Question-Answer Pairs for TV Series","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Paragraph; Computer science; Information retrieval; Series (stratigraphy); Snapshot (computer storage); Set (abstract data type); Television series; Questions and answers; Pipeline (software); Artificial intelligence; Natural language processing; Data science; World Wide Web; Database; Media studies","score_opus":0.020473852071755335,"score_gpt":0.2545611907355675,"score_spread":0.23408733866381215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225656840","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01633671,0.0008234943,0.70570046,0.0008355875,0.00079099723,0.0011030143,0.015564081,0.21296096,0.045884747],"genre_scores_gemma":[0.22730784,0.0005795187,0.64540803,0.0009763281,0.00029737956,0.0016519615,0.038932238,0.020796336,0.06405037],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982687,0.00052960915,0.00010876522,0.00044598646,0.00048692685,0.00015990043],"domain_scores_gemma":[0.9981839,0.0011343483,0.000047587695,0.00025982904,0.00022927758,0.0001449895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013437811,0.0022199482,0.0014065244,0.0012089227,0.00084060937,0.0021699013,0.002846482,0.0018050737,0.08407521],"category_scores_gemma":[0.0069671962,0.0009769185,0.0014228473,0.0007192731,0.00064228964,0.005074841,0.004397388,0.0019922005,0.026279218],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002358342,0.00072473934,0.0013180516,0.0019213088,0.00017146564,0.00044278754,0.0018320866,0.009006026,0.033139605,0.07754873,0.44068295,0.4308538],"study_design_scores_gemma":[0.00079444674,0.0005869646,0.0014126265,0.0002488806,0.00016768227,0.0004771838,0.00076534314,0.26609567,0.03960651,0.14871053,0.5409087,0.00022551705],"about_ca_topic_score_codex":0.0038776225,"about_ca_topic_score_gemma":0.006311394,"teacher_disagreement_score":0.08407521,"about_ca_system_score_codex":0.0008983454,"about_ca_system_score_gemma":0.0008981764,"threshold_uncertainty_score":0.2812596},"labels":[],"label_agreement":null},{"id":"W4225661091","doi":"10.3758/s13421-022-01300-7","title":"Sources and destinations of misattributions in recall of instances of repeated events","year":2022,"lang":"en","type":"article","venue":"Memory & Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University","funders":"Economic and Social Research Council; University of Portsmouth","keywords":"Psychology; Schema (genetic algorithms); Recall; Attribution; Confusion; Inference; Anagrams; Context (archaeology); Social psychology; Repeated measures design; Cognitive psychology; Natural language processing; Information retrieval; Artificial intelligence; Computer science; Statistics; Mathematics","score_opus":0.02982022865342573,"score_gpt":0.25005393387434627,"score_spread":0.22023370522092053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225661091","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98485214,0.0006378345,0.011027227,0.00018136394,0.000051877363,0.000043566215,0.00039605517,0.00023987622,0.0025700785],"genre_scores_gemma":[0.99644166,0.00015903436,0.0023026597,0.00002093337,0.000030370566,0.000013205043,0.00035943688,0.000087776454,0.00058486586],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9982127,0.0005085175,0.0002681074,0.00042164294,0.0004404278,0.00014868217],"domain_scores_gemma":[0.9286154,0.05279437,0.00640414,0.0068250042,0.0043172287,0.0010437631],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035414584,0.00042134448,0.00056890654,0.002284223,0.00048332612,0.003687767,0.00094849744,0.0015780149,0.0021424661],"category_scores_gemma":[0.07501151,0.00058936974,0.0005869846,0.001622078,0.00058631756,0.003432056,0.001525474,0.0017095837,0.00049292744],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011352724,0.00074504083,0.6797424,0.00077933655,0.0010907806,0.002299624,0.02334976,0.012675397,0.047645614,0.016261797,0.00212553,0.20193203],"study_design_scores_gemma":[0.00047738018,0.0012317699,0.728886,0.00043777341,0.001816063,0.006946235,0.011267822,0.1257875,0.050519504,0.066880815,0.005361618,0.0003875596],"about_ca_topic_score_codex":0.0023914773,"about_ca_topic_score_gemma":0.0012409136,"teacher_disagreement_score":0.003687767,"about_ca_system_score_codex":0.00056263857,"about_ca_system_score_gemma":0.00055107125,"threshold_uncertainty_score":0.01872921},"labels":[],"label_agreement":null},{"id":"W4226003933","doi":"10.18653/v1/2022.acl-long.96","title":"Better Language Model with Hypernym Class Prediction","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Microsoft Research","keywords":"Perplexity; Computer science; Artificial intelligence; Transformer; Class (philosophy); Language model; WordNet; Security token; Natural language processing; Machine learning; Context (archaeology)","score_opus":0.006999110390969953,"score_gpt":0.20988021284321878,"score_spread":0.20288110245224883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226003933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28786042,0.0011096196,0.6947632,0.0022147356,0.0003343703,0.00015316182,0.001551862,0.007880093,0.0041325996],"genre_scores_gemma":[0.8723606,0.00023447724,0.11871059,0.00043844857,0.00017015864,0.00015306941,0.0023116088,0.00037592443,0.0052450933],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931216,0.00029644318,0.00003567539,0.00024221977,0.00006409108,0.000049478967],"domain_scores_gemma":[0.9977266,0.0014227136,0.00008097592,0.00038844653,0.00030354,0.00007765708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022809869,0.0008892023,0.00081057113,0.0010381023,0.00041088025,0.001415289,0.001327165,0.0009640639,0.0023123259],"category_scores_gemma":[0.0065876,0.00031778528,0.0011214847,0.00086322037,0.0003264367,0.0049447496,0.0009765923,0.0027315512,0.0017869916],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006453694,0.00045395867,0.010439646,0.00021360094,0.00029691306,0.00017927149,0.0005545094,0.5271757,0.0125651425,0.015247999,0.014331148,0.41789672],"study_design_scores_gemma":[0.000008577635,0.000015671192,0.0001406547,0.000004270629,0.000009671394,0.0000081885055,0.000012778322,0.9947524,0.0007052262,0.0038970907,0.00044065207,0.000004695609],"about_ca_topic_score_codex":0.0058812504,"about_ca_topic_score_gemma":0.009560985,"teacher_disagreement_score":0.0058812504,"about_ca_system_score_codex":0.0007298545,"about_ca_system_score_gemma":0.0010248091,"threshold_uncertainty_score":0.012063146},"labels":[],"label_agreement":null},{"id":"W4226059645","doi":"10.1162/tacl_a_00471","title":"TopiOCQA: Open-domain Conversational Question Answering with Topic Switching","year":2022,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Minnow Environmental (Canada); Research Canada; Microsoft (Canada); McGill University","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Conversation; Computer science; Question answering; Open domain; Domain (mathematical analysis); Information retrieval; Interdependence; Natural language processing; Artificial intelligence; Relevance (law); Code (set theory); Linguistics; Set (abstract data type)","score_opus":0.015187584352610738,"score_gpt":0.25792401609927124,"score_spread":0.24273643174666049,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226059645","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21077205,0.012377717,0.21660125,0.0051048356,0.0017385945,0.0050865244,0.3328103,0.18974109,0.025767628],"genre_scores_gemma":[0.291302,0.00088209467,0.20516017,0.0017817171,0.00039242097,0.003031041,0.48823872,0.0015715217,0.0076402943],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99447244,0.0023815425,0.0003820398,0.0016350816,0.00076294184,0.0003658794],"domain_scores_gemma":[0.9894555,0.005302078,0.00042679417,0.002410132,0.0016174435,0.0007880501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041930336,0.0024624004,0.0013719068,0.003432803,0.001971981,0.0026590226,0.0044109654,0.0030631414,0.007699486],"category_scores_gemma":[0.020064658,0.0006903227,0.0019215929,0.0024376323,0.0009859382,0.005132341,0.005119776,0.0035130072,0.006802582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00328068,0.0018235833,0.01207427,0.0049089156,0.00063587487,0.0010127898,0.0030712718,0.03275208,0.02015776,0.008756197,0.6344944,0.27703208],"study_design_scores_gemma":[0.0010290723,0.0008591131,0.015385165,0.00048760648,0.0003217899,0.0011904399,0.0022608284,0.6801566,0.023591295,0.02129519,0.25304788,0.00037505914],"about_ca_topic_score_codex":0.039864566,"about_ca_topic_score_gemma":0.045868874,"teacher_disagreement_score":0.039864566,"about_ca_system_score_codex":0.0024400011,"about_ca_system_score_gemma":0.0031037575,"threshold_uncertainty_score":0.07926506},"labels":[],"label_agreement":null},{"id":"W4226145103","doi":"10.18653/v1/2022.findings-acl.326","title":"On the data requirements of probing","year":2022,"lang":"en","type":"article","venue":"Findings of the Association for Computational Linguistics: ACL 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Reliability (semiconductor); Context (archaeology); Construct (python library); Data mining; Machine learning; Artificial intelligence; Power (physics)","score_opus":0.054162738601187525,"score_gpt":0.291028757242492,"score_spread":0.23686601864130447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226145103","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12935854,0.007681459,0.6953404,0.11801296,0.0022876658,0.0019412522,0.012944842,0.006339601,0.026093261],"genre_scores_gemma":[0.50664765,0.0023372394,0.44610903,0.013650176,0.0023368895,0.0041142846,0.01786789,0.0027051382,0.004231726],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.7993628,0.14663368,0.010859938,0.017439358,0.022870181,0.0028341473],"domain_scores_gemma":[0.14934397,0.7665843,0.006776336,0.059332296,0.014480481,0.0034826177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.15146323,0.0020775236,0.005660926,0.003211724,0.004348049,0.011224856,0.01005458,0.0107450625,0.011888778],"category_scores_gemma":[0.685493,0.0037523624,0.0031971177,0.007854061,0.010096867,0.034245517,0.014305756,0.014496589,0.0047530113],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009351456,0.0011847921,0.054059166,0.0027866417,0.0007923673,0.0015822558,0.006661666,0.066939525,0.008518392,0.44857207,0.12005045,0.27950126],"study_design_scores_gemma":[0.0008555682,0.00048458087,0.0069303177,0.00051169546,0.00019131607,0.0012455896,0.0024988814,0.23806246,0.003625033,0.719126,0.02631887,0.00014966173],"about_ca_topic_score_codex":0.0046378314,"about_ca_topic_score_gemma":0.0035095837,"teacher_disagreement_score":0.15146323,"about_ca_system_score_codex":0.0037036038,"about_ca_system_score_gemma":0.006071488,"threshold_uncertainty_score":0.80102366},"labels":[],"label_agreement":null},{"id":"W4226208814","doi":"10.1007/978-3-030-99736-6_41","title":"Another Look at DPR: Reproduction of Training and Replication of Retrieval","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Question answering; Replication (statistics); Codebase; Ranking (information retrieval); Scope (computer science); Encoder; Artificial intelligence; Natural language processing; Programming language; Source code","score_opus":0.03770741802576652,"score_gpt":0.2571676180958307,"score_spread":0.2194602000700642,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226208814","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013053104,0.0018366496,0.6409673,0.011226422,0.0027509644,0.00025197206,0.0013038642,0.023214372,0.3053953],"genre_scores_gemma":[0.19160445,0.0011576179,0.505405,0.004472386,0.0014605719,0.00032435634,0.0015480906,0.014753218,0.27927428],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985044,0.0005468621,0.000079172954,0.0003413844,0.0004352939,0.00009296003],"domain_scores_gemma":[0.9892285,0.0037450227,0.00013352743,0.0059313974,0.0007227586,0.00023887455],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0023021705,0.00061491044,0.0006874272,0.0010097086,0.0009684897,0.0031459269,0.0028663925,0.0014205404,0.061511],"category_scores_gemma":[0.015428376,0.0006546184,0.000851207,0.0012339231,0.0022207059,0.009268157,0.002721284,0.0033229932,0.018686833],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002999476,0.00020463625,0.0003251575,0.00036632954,0.000024198302,0.00021495124,0.0016507545,0.0039356723,0.023952255,0.28680533,0.114175506,0.5680453],"study_design_scores_gemma":[0.00016172059,0.00027184046,0.0010158687,0.0001615384,0.000057824713,0.0011840677,0.0006353341,0.04220979,0.031356547,0.2208529,0.701957,0.00013563645],"about_ca_topic_score_codex":0.0025667239,"about_ca_topic_score_gemma":0.0026426844,"teacher_disagreement_score":0.99769783,"about_ca_system_score_codex":0.0010838779,"about_ca_system_score_gemma":0.0009895312,"threshold_uncertainty_score":0.20577478},"labels":[],"label_agreement":null},{"id":"W4226218094","doi":"10.1145/3486622.3494010","title":"Relation Extraction with Sentence Simplification Process and Entity Information","year":2021,"lang":"en","type":"article","venue":"IEEE/WIC/ACM International Conference on Web Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Process (computing); Sentence; Relation (database); Relationship extraction; Information extraction; Artificial intelligence; Information retrieval; Programming language; Database","score_opus":0.05525394880722037,"score_gpt":0.3218998359069743,"score_spread":0.2666458870997539,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226218094","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028463058,0.0012451757,0.9553074,0.00043129575,0.00017002798,0.00048637408,0.0023310583,0.009113824,0.0024518135],"genre_scores_gemma":[0.14740641,0.0010098944,0.83090556,0.00030656165,0.00017304713,0.0003686122,0.014149376,0.00034573578,0.0053349],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989312,0.00020818616,0.00014917494,0.00039580424,0.00026018207,0.000055309603],"domain_scores_gemma":[0.9982889,0.00071678543,0.00018402992,0.00036830755,0.00040208636,0.00003982265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011041078,0.0015648856,0.0011820203,0.0029643702,0.000729721,0.0009780369,0.0018090844,0.0008492869,0.002736687],"category_scores_gemma":[0.0042539365,0.00047597048,0.0014786039,0.002965041,0.0005151871,0.00315424,0.0013124194,0.001556133,0.0020270862],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034735654,0.00026523139,0.0021659997,0.00081862987,0.00021456674,0.00083452114,0.00062918337,0.03031429,0.048889697,0.0065064267,0.018419977,0.8905942],"study_design_scores_gemma":[0.00008424648,0.00026287048,0.0037336973,0.00010421501,0.00038650254,0.0009108465,0.00033889723,0.85405886,0.08192243,0.016167639,0.041930277,0.00009951105],"about_ca_topic_score_codex":0.006170884,"about_ca_topic_score_gemma":0.010479845,"teacher_disagreement_score":0.006170884,"about_ca_system_score_codex":0.00060289056,"about_ca_system_score_gemma":0.0014472449,"threshold_uncertainty_score":0.012269914},"labels":[],"label_agreement":null},{"id":"W4226316739","doi":"10.2196/34834","title":"Pretrained Transformer Language Models Versus Pretrained Word Embeddings for the Detection of Accurate Health Information on Arabic Social Media: Comparative Study","year":2022,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Taibah University; Science Foundation Ireland","keywords":"Computer science; Social media; Natural language processing; Leverage (statistics); Language model; Artificial intelligence; Transformer; Arabic; Health informatics; Machine learning; Information retrieval; World Wide Web; Linguistics; Medicine; Public health","score_opus":0.14442572129860515,"score_gpt":0.4327990415068596,"score_spread":0.28837332020825446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226316739","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9341162,0.0063155205,0.047747694,0.00092307216,0.0006409636,0.0002487116,0.0019029805,0.0023412195,0.005763644],"genre_scores_gemma":[0.9662161,0.0012051472,0.024575595,0.00030964674,0.00009656853,0.00013241312,0.004847942,0.00013353338,0.0024831463],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985689,0.0005744651,0.00013504605,0.00037196843,0.00018466212,0.00016501402],"domain_scores_gemma":[0.9937354,0.0042186445,0.00024939768,0.00047019872,0.0011621161,0.00016428829],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030359772,0.0027781073,0.00090730673,0.001119575,0.00040373797,0.0012377051,0.0012320142,0.0013624901,0.0023248186],"category_scores_gemma":[0.011332,0.0005526656,0.00097217056,0.00074544444,0.0006451747,0.003043144,0.0012503678,0.002902193,0.0019065511],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028925156,0.0023351794,0.04981293,0.0013446728,0.0009617928,0.00077856606,0.00082123675,0.33109114,0.0099249,0.0015218542,0.015869673,0.58264554],"study_design_scores_gemma":[0.00005430739,0.0005556446,0.00552586,0.00009201175,0.00015803907,0.00013642672,0.00040715223,0.9841897,0.0065195854,0.0010272962,0.0012844002,0.000049584058],"about_ca_topic_score_codex":0.015160346,"about_ca_topic_score_gemma":0.013453497,"teacher_disagreement_score":0.015160346,"about_ca_system_score_codex":0.0011683216,"about_ca_system_score_gemma":0.0011346132,"threshold_uncertainty_score":0.030144215},"labels":[],"label_agreement":null},{"id":"W4226334043","doi":"10.1016/j.ipm.2022.102933","title":"ARL: An adaptive reinforcement learning framework for complex question answering over knowledge base","year":2022,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":55,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Interpretability; Computer science; Reinforcement learning; Artificial intelligence; Knowledge base; Benchmark (surveying); Metric (unit); Machine learning; Relation (database); Question answering; Data mining","score_opus":0.03806807387468865,"score_gpt":0.30122174349742853,"score_spread":0.2631536696227399,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226334043","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003173174,0.00014107133,0.9893959,0.00015311614,0.000038873797,0.00008921047,0.00018235784,0.0059997295,0.00082668563],"genre_scores_gemma":[0.20211273,0.00023232646,0.79039913,0.00033709523,0.00009015576,0.00045210437,0.0008121919,0.0006992914,0.0048649493],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991291,0.00034367715,0.00004634295,0.00019153851,0.00020579835,0.000083528124],"domain_scores_gemma":[0.99828476,0.001060062,0.00008760417,0.00020119397,0.0002476737,0.00011876174],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022539534,0.0007471623,0.0010364782,0.00070881407,0.00041542447,0.0012235763,0.00340411,0.0014098046,0.00848664],"category_scores_gemma":[0.005932056,0.0006166711,0.0008977554,0.0005369014,0.00070722937,0.0020518748,0.002092668,0.0025045052,0.0022313762],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051587285,0.00052602997,0.0011586137,0.0003918624,0.00017255056,0.00020070329,0.0003509349,0.45002756,0.007402636,0.032241948,0.017626304,0.489385],"study_design_scores_gemma":[0.000026997126,0.000023755658,0.000049229864,0.000007105842,0.000010103513,0.0000099509525,0.000008611914,0.9889051,0.00064854336,0.008557246,0.0017464084,0.000006968981],"about_ca_topic_score_codex":0.008216849,"about_ca_topic_score_gemma":0.011650161,"teacher_disagreement_score":0.00848664,"about_ca_system_score_codex":0.0009485146,"about_ca_system_score_gemma":0.0012721429,"threshold_uncertainty_score":0.028390586},"labels":[],"label_agreement":null},{"id":"W4226464466","doi":"10.2139/ssrn.4068677","title":"A Study of Update Request Comments in Stack Overflow Answer Posts","year":2022,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Stack (abstract data type); Computer science; Computer security; Programming language","score_opus":0.014036207997679627,"score_gpt":0.26082692891346104,"score_spread":0.24679072091578141,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226464466","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99300295,0.00015392722,0.0011729475,0.00029708954,0.00003498002,0.00009769888,0.00064888527,0.00009828342,0.0044931564],"genre_scores_gemma":[0.99418324,0.000121037934,0.0009678576,0.000062069244,0.000069252696,0.00007314235,0.0012307417,0.00006196025,0.0032307827],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99582714,0.0019182211,0.00027647588,0.00032785517,0.0012692218,0.00038105532],"domain_scores_gemma":[0.80792505,0.15706511,0.016133273,0.0024803467,0.013964464,0.00243173],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002999089,0.0002934313,0.00031999368,0.0036296197,0.0017525193,0.0028208834,0.0005902402,0.0012047539,0.0048038512],"category_scores_gemma":[0.07377495,0.00028786503,0.00029165522,0.004628538,0.00067950704,0.0027462349,0.0007270932,0.0012578706,0.0015233794],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023074267,0.0014811916,0.8318593,0.0007027127,0.0001222503,0.0014239931,0.053651746,0.0016753828,0.010170275,0.0053539667,0.01025254,0.0809993],"study_design_scores_gemma":[0.00006983338,0.0009393852,0.8855004,0.0002694003,0.00019721367,0.0010070879,0.056966685,0.025466807,0.0052402695,0.0016022346,0.022637814,0.0001029175],"about_ca_topic_score_codex":0.013635874,"about_ca_topic_score_gemma":0.015343915,"teacher_disagreement_score":0.013635874,"about_ca_system_score_codex":0.0015533898,"about_ca_system_score_gemma":0.00142217,"threshold_uncertainty_score":0.02711302},"labels":[],"label_agreement":null},{"id":"W4229005744","doi":"10.18653/v1/2022.naacl-main.257","title":"Learning to Transfer Prompts for Text Generation","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Renmin University of China; National Natural Science Foundation of China","keywords":"Computational linguistics; Computer science; Natural language processing; Linguistics; Artificial intelligence; Association (psychology); Cognitive science; Philosophy; Psychology; Epistemology","score_opus":0.02267813460258396,"score_gpt":0.25319064266185753,"score_spread":0.23051250805927356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229005744","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036699694,0.0012804648,0.8896211,0.0015036619,0.0018329973,0.00069187314,0.0020302301,0.057879534,0.008460503],"genre_scores_gemma":[0.53716725,0.00057251344,0.43894,0.00079775334,0.0007814942,0.0012051995,0.0057690283,0.0017118194,0.013054955],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979327,0.00083618914,0.00010934065,0.0007156426,0.00023492973,0.00017123544],"domain_scores_gemma":[0.99240667,0.00482651,0.00023474461,0.0012242862,0.00095444487,0.0003533685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029312451,0.0016782256,0.0010495101,0.0014058991,0.00090195105,0.0012558427,0.0021711425,0.0018195334,0.021106357],"category_scores_gemma":[0.01769433,0.00060774304,0.0007820611,0.0009788381,0.00073201564,0.005089655,0.003498132,0.00292162,0.011959233],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011375227,0.00044113258,0.0012802106,0.000308137,0.000040759296,0.000181334,0.00035559692,0.013132505,0.008147239,0.009661951,0.05268065,0.91263306],"study_design_scores_gemma":[0.00047000352,0.00047120135,0.0007085947,0.00008063581,0.000076344375,0.00018565888,0.00032071682,0.8569023,0.01948575,0.09998628,0.021258164,0.000054325556],"about_ca_topic_score_codex":0.0012023774,"about_ca_topic_score_gemma":0.0018329115,"teacher_disagreement_score":0.021106357,"about_ca_system_score_codex":0.00071944454,"about_ca_system_score_gemma":0.0014254175,"threshold_uncertainty_score":0.07060784},"labels":[],"label_agreement":null},{"id":"W4229451765","doi":"10.3390/jimaging8050131","title":"BI-RADS BERT and Using Section Segmentation to Understand Radiology Reports","year":2022,"lang":"en","type":"article","venue":"Journal of Imaging","topic":"Topic Modeling","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sunnybrook Health Science Centre; University of Toronto","funders":"Simon Fraser University; Compute Canada; Canadian Institutes of Health Research; Sunnybrook Research Institute","keywords":"Computer science; Segmentation; Artificial intelligence; Lexicon; Sentence; Natural language processing; Breast imaging; Section (typography); Classifier (UML); Mammography; Breast cancer; Pattern recognition (psychology); Medicine; Cancer","score_opus":0.02553176897951776,"score_gpt":0.27781596360131805,"score_spread":0.2522841946218003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229451765","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19898209,0.0012823149,0.7402887,0.0013615452,0.00043153006,0.0005411412,0.0069009555,0.039289705,0.010922012],"genre_scores_gemma":[0.55368984,0.00068418484,0.41319478,0.00041709005,0.00013899303,0.00032601666,0.01983872,0.0013011176,0.010409343],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992749,0.00020313417,0.00008133875,0.00022185805,0.00015300287,0.00006581287],"domain_scores_gemma":[0.9982008,0.00082647224,0.0002098414,0.00028897086,0.00041884033,0.00005514154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001106554,0.0011042993,0.0003091148,0.0013217039,0.0002591286,0.0011669747,0.0007200975,0.0008678995,0.003072529],"category_scores_gemma":[0.004202924,0.00046505878,0.0007716601,0.0004786722,0.00037192027,0.0037238467,0.0010305189,0.0010481932,0.0038031084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010328499,0.00035958758,0.022819342,0.0008567336,0.00014794503,0.00084426964,0.0016501338,0.100390315,0.09869079,0.013845493,0.036764253,0.72259825],"study_design_scores_gemma":[0.000034833527,0.00030301278,0.010470688,0.00012825963,0.00009741551,0.0007918689,0.0004678168,0.8735829,0.055806406,0.008443163,0.049777076,0.00009655958],"about_ca_topic_score_codex":0.007635999,"about_ca_topic_score_gemma":0.009459048,"teacher_disagreement_score":0.007635999,"about_ca_system_score_codex":0.00080884586,"about_ca_system_score_gemma":0.0009612197,"threshold_uncertainty_score":0.015183151},"labels":[],"label_agreement":null},{"id":"W4230743089","doi":"10.4018/978-1-60566-050-9.ch048","title":"Information Retrieval by Semantic Similarity","year":2011,"lang":"en","type":"book-chapter","venue":"Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Semantic similarity; WordNet; Information retrieval; Computer science; Explicit semantic analysis; Ontology; Similarity (geometry); Semantic integration; Semantic computing; Vector space model; Semantic search; Natural language processing; Ontology-based data integration; Artificial intelligence; Semantic technology; Semantic Web","score_opus":0.02606271963612654,"score_gpt":0.2341778444848152,"score_spread":0.20811512484868866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230743089","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0086867325,0.021449132,0.9299759,0.0025456455,0.0007796675,0.0010795405,0.0013878638,0.001911513,0.032183886],"genre_scores_gemma":[0.16155598,0.026421668,0.7823173,0.0012822098,0.0017175035,0.0015832573,0.0057847737,0.00040349178,0.018933872],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9909503,0.0036954575,0.0008152655,0.0010852567,0.0031953047,0.00025846373],"domain_scores_gemma":[0.99662614,0.0016967695,0.00023975055,0.0008403769,0.00052802573,0.000069040056],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004831193,0.0014816497,0.0027306916,0.013834187,0.0011058105,0.007183242,0.0023415235,0.0025965367,0.010131286],"category_scores_gemma":[0.015346466,0.0006111821,0.0019273389,0.019284718,0.0023593227,0.015415291,0.004669832,0.0018510706,0.0083139185],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014640813,0.00021624619,0.0010496494,0.0020186324,0.00034754307,0.0002448589,0.0005220021,0.012049104,0.006846529,0.25633004,0.043581136,0.67664784],"study_design_scores_gemma":[0.000109355904,0.0003098661,0.0014283315,0.000641543,0.0002060098,0.0013755177,0.0006508791,0.15048249,0.00974799,0.64387906,0.19096166,0.00020732661],"about_ca_topic_score_codex":0.0014514558,"about_ca_topic_score_gemma":0.0011087265,"teacher_disagreement_score":0.013834187,"about_ca_system_score_codex":0.0023511471,"about_ca_system_score_gemma":0.0016260087,"threshold_uncertainty_score":0.033892572},"labels":[],"label_agreement":null},{"id":"W4231122779","doi":"10.18653/v1/2021.emnlp-main.288","title":"Automated Generation of Accurate &amp; Fluent Medical X-ray Reports","year":2021,"lang":"en","type":"article","venue":"Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Fluency; Embedding; Generator (circuit theory); Transformer; Natural language processing; Artificial intelligence; Information retrieval; Linguistics","score_opus":0.09840819144765348,"score_gpt":0.4311545131288209,"score_spread":0.33274632168116747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231122779","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045215603,0.0006484048,0.91022784,0.0006249353,0.00023252034,0.00040037834,0.0037484374,0.03680521,0.0020966865],"genre_scores_gemma":[0.28485164,0.00054102193,0.69478774,0.00032869107,0.00020606657,0.00041421544,0.012507989,0.0026614699,0.0037012077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99838376,0.00059277157,0.00014191908,0.00039544926,0.0004050018,0.00008110123],"domain_scores_gemma":[0.9920467,0.005440183,0.0005026309,0.0009599765,0.00087776827,0.0001728486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021569107,0.0015025832,0.00073806057,0.0014880182,0.0001871821,0.0011448392,0.0017043175,0.0010133171,0.005209121],"category_scores_gemma":[0.01298163,0.00052934623,0.0011580989,0.00062255765,0.00042654687,0.0012692036,0.0016917702,0.0008990154,0.003587427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012124905,0.00037557498,0.006043965,0.0015729514,0.0002024853,0.00293525,0.00089531596,0.09937468,0.063964985,0.009044558,0.04680929,0.7675684],"study_design_scores_gemma":[0.00023567337,0.00031862396,0.0021509847,0.00013216535,0.00014425149,0.0019218021,0.0001877435,0.83553445,0.11786331,0.017765932,0.023665497,0.000079565754],"about_ca_topic_score_codex":0.0006545288,"about_ca_topic_score_gemma":0.0008831666,"teacher_disagreement_score":0.005209121,"about_ca_system_score_codex":0.0004288535,"about_ca_system_score_gemma":0.0008585897,"threshold_uncertainty_score":0.017426252},"labels":[],"label_agreement":null},{"id":"W4231235349","doi":"10.32920/ryerson.14661732.v1","title":"Benchmarking of semantic annotation systems","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Benchmarking; Annotation; Natural language processing; Information retrieval; Semantic annotation; Artificial intelligence; Data science","score_opus":0.03099480714304642,"score_gpt":0.25698546357188523,"score_spread":0.2259906564288388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231235349","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4622283,0.007469304,0.41080973,0.002371416,0.002237285,0.0033767838,0.010026863,0.023246352,0.078233965],"genre_scores_gemma":[0.6669275,0.0019596245,0.26553074,0.0005701125,0.00033465665,0.0019652615,0.047943216,0.003935444,0.010833434],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.8991154,0.057112437,0.008420585,0.010615142,0.022504682,0.0022317937],"domain_scores_gemma":[0.8825113,0.047629025,0.004100754,0.027266538,0.035919975,0.0025724112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06453384,0.0018321422,0.0014883144,0.009161702,0.0032088875,0.007497363,0.003660603,0.0030693712,0.006315471],"category_scores_gemma":[0.120766744,0.0007116192,0.0015037812,0.0063498,0.0019679507,0.010806801,0.0067696907,0.0018955718,0.005438419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032204154,0.003286914,0.042015366,0.00786321,0.0019645046,0.0008581761,0.01197138,0.049825028,0.045315895,0.028660094,0.048450265,0.75656873],"study_design_scores_gemma":[0.0005886261,0.005346079,0.103142306,0.0032212017,0.0013768246,0.0023105831,0.017990783,0.35933322,0.15205234,0.047579788,0.30611113,0.00094715704],"about_ca_topic_score_codex":0.0032667753,"about_ca_topic_score_gemma":0.004176205,"teacher_disagreement_score":0.06453384,"about_ca_system_score_codex":0.003113991,"about_ca_system_score_gemma":0.0036893792,"threshold_uncertainty_score":0.3412916},"labels":[],"label_agreement":null},{"id":"W4231451126","doi":"10.3115/1642025.1642031","title":"A cognitive model for the representation and acquisition of verb selectional preferences","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Verb; Computer science; Argument (complex analysis); Alternation (linguistics); Natural language processing; Artificial intelligence; Object (grammar); Reflexive verb; Representation (politics); Set (abstract data type); Cognition; Linguistics; Modal verb; Psychology; Programming language","score_opus":0.07026216968939814,"score_gpt":0.3268636723456144,"score_spread":0.2566015026562163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231451126","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.078761466,0.0001714573,0.91058916,0.0012564643,0.00003581677,0.00012067349,0.00032915676,0.00040701556,0.008328797],"genre_scores_gemma":[0.8004723,0.00022034506,0.1953635,0.00035808704,0.00005379736,0.00039757387,0.0006451713,0.00008957249,0.0023995736],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985715,0.0005268456,0.00007202763,0.00042834654,0.00027133134,0.00012993331],"domain_scores_gemma":[0.9928752,0.0046921014,0.00059508224,0.00092543615,0.0006149449,0.00029717488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029623895,0.0005661851,0.0006118459,0.0013225399,0.00065039855,0.002872591,0.0021545526,0.0011944927,0.0046468233],"category_scores_gemma":[0.015267078,0.00071914296,0.0016863622,0.001039517,0.0020345973,0.0050247433,0.0013579392,0.0024668644,0.0005325665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048697848,0.00046194548,0.011927885,0.0004728459,0.00046750822,0.00055279647,0.004734862,0.107599355,0.021955973,0.6954881,0.0042294757,0.15162219],"study_design_scores_gemma":[0.000070293456,0.0001211597,0.0034258119,0.0000326507,0.00009148335,0.00032071883,0.0001911239,0.3781871,0.0022325562,0.6132339,0.0020319119,0.00006130268],"about_ca_topic_score_codex":0.0038149995,"about_ca_topic_score_gemma":0.003506643,"teacher_disagreement_score":0.0046468233,"about_ca_system_score_codex":0.0014874345,"about_ca_system_score_gemma":0.0011437316,"threshold_uncertainty_score":0.015666842},"labels":[],"label_agreement":null},{"id":"W4233906699","doi":"10.1007/10985687_6","title":"Neural Probabilistic Language Models","year":2006,"lang":"en","type":"book-chapter","venue":"Studies in fuzziness and soft computing","topic":"Topic Modeling","field":"Computer Science","cited_by":132,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal","funders":"","keywords":"Computer science; Generalization; Sentence; Curse of dimensionality; Artificial intelligence; Language model; Word (group theory); Natural language processing; Probabilistic logic; Representation (politics); Sequence (biology); Set (abstract data type); Linguistics; Mathematics","score_opus":0.056090077404893204,"score_gpt":0.287402023592984,"score_spread":0.2313119461880908,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233906699","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011490583,0.0056890706,0.9467938,0.0027756859,0.00036457943,0.00002722255,0.00043309043,0.0012119767,0.031214025],"genre_scores_gemma":[0.64411217,0.009672135,0.25746024,0.00087563234,0.0011110181,0.00026797937,0.00247419,0.0006134571,0.0834133],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995616,0.00014035603,0.00002137301,0.00011364481,0.0001331687,0.00002977463],"domain_scores_gemma":[0.9984333,0.0011409594,0.00006241853,0.00016939337,0.00015980162,0.00003416432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008027947,0.0006190033,0.0008226144,0.00080580945,0.00046107575,0.0019058926,0.0013140924,0.0010974476,0.0067877136],"category_scores_gemma":[0.0052618007,0.0005861519,0.000715029,0.0011666501,0.00090206845,0.0036983308,0.0008058605,0.002361336,0.0022041826],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005887798,0.000061835075,0.00044558127,0.00017145903,0.00009541112,0.00008216203,0.00022197867,0.123059236,0.0020324676,0.63067245,0.01820091,0.22489768],"study_design_scores_gemma":[0.00000525985,0.000009304747,0.0001823774,0.000022873845,0.000019840625,0.000064062086,0.00002005938,0.48232594,0.0005989685,0.5086477,0.008087834,0.00001582531],"about_ca_topic_score_codex":0.0019029549,"about_ca_topic_score_gemma":0.0024362842,"teacher_disagreement_score":0.0067877136,"about_ca_system_score_codex":0.0007533101,"about_ca_system_score_gemma":0.0006154235,"threshold_uncertainty_score":0.022707164},"labels":[],"label_agreement":null},{"id":"W4235795010","doi":"10.32920/ryerson.14661732","title":"Benchmarking of semantic annotation systems","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Benchmarking; Annotation; Natural language processing; Information retrieval; Semantic annotation; Artificial intelligence","score_opus":0.03099480714304642,"score_gpt":0.25698546357188523,"score_spread":0.2259906564288388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4235795010","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4622283,0.007469304,0.41080973,0.002371416,0.002237285,0.0033767838,0.010026863,0.023246352,0.078233965],"genre_scores_gemma":[0.6669275,0.0019596245,0.26553074,0.0005701125,0.00033465665,0.0019652615,0.047943216,0.003935444,0.010833434],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.8991154,0.057112437,0.008420585,0.010615142,0.022504682,0.0022317937],"domain_scores_gemma":[0.8825113,0.047629025,0.004100754,0.027266538,0.035919975,0.0025724112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06453384,0.0018321422,0.0014883144,0.009161702,0.0032088875,0.007497363,0.003660603,0.0030693712,0.006315471],"category_scores_gemma":[0.120766744,0.0007116192,0.0015037812,0.0063498,0.0019679507,0.010806801,0.0067696907,0.0018955718,0.005438419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032204154,0.003286914,0.042015366,0.00786321,0.0019645046,0.0008581761,0.01197138,0.049825028,0.045315895,0.028660094,0.048450265,0.75656873],"study_design_scores_gemma":[0.0005886261,0.005346079,0.103142306,0.0032212017,0.0013768246,0.0023105831,0.017990783,0.35933322,0.15205234,0.047579788,0.30611113,0.00094715704],"about_ca_topic_score_codex":0.0032667753,"about_ca_topic_score_gemma":0.004176205,"teacher_disagreement_score":0.06453384,"about_ca_system_score_codex":0.003113991,"about_ca_system_score_gemma":0.0036893792,"threshold_uncertainty_score":0.3412916},"labels":[],"label_agreement":null},{"id":"W4236183820","doi":"10.32920/ryerson.14654349.v1","title":"Identifying User Interests In An Online Discussion Forum With Deep Learning","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Ontario Tech University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Metric (unit); Artificial neural network; Set (abstract data type); Sample (material); Artificial intelligence; Probabilistic logic; Machine learning; Recommender system; Social media; Test set; Data set; Data mining; Information retrieval; World Wide Web; Engineering","score_opus":0.05758956287715514,"score_gpt":0.3124403340316915,"score_spread":0.25485077115453636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236183820","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8202277,0.00028752454,0.17261006,0.0004896334,0.00004781389,0.00021759386,0.0008688229,0.001290033,0.003960858],"genre_scores_gemma":[0.9647848,0.000042518124,0.032620087,0.00003559773,0.000012849975,0.00008693241,0.0005020997,0.00002202826,0.0018929844],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987405,0.0005914906,0.000046782523,0.00029446435,0.00016820095,0.00015854013],"domain_scores_gemma":[0.99526274,0.003257358,0.00032928138,0.00030609677,0.00059361954,0.0002510077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003897632,0.0006347803,0.00043422714,0.001758392,0.00061972876,0.0013069386,0.00078715984,0.0009584333,0.0013162888],"category_scores_gemma":[0.00879554,0.0004353188,0.00058346335,0.0009860641,0.00040621168,0.0022195915,0.001420602,0.0016005208,0.00061366806],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025703865,0.002870197,0.15296829,0.0005094622,0.0003772162,0.00053403754,0.00891769,0.26738402,0.038046665,0.01428694,0.0076489584,0.50388616],"study_design_scores_gemma":[0.00001225359,0.00008140787,0.006075306,0.000011155521,0.000012881516,0.000022224156,0.0001931779,0.9864362,0.0028380991,0.0035877344,0.00071511074,0.000014505789],"about_ca_topic_score_codex":0.00680641,"about_ca_topic_score_gemma":0.011204113,"teacher_disagreement_score":0.00680641,"about_ca_system_score_codex":0.0013075715,"about_ca_system_score_gemma":0.0007580225,"threshold_uncertainty_score":0.020612895},"labels":[],"label_agreement":null},{"id":"W4238310921","doi":"10.1109/icosc.2007.4338410","title":"Comparing the Contribution of Syntactic and Semantic Features in Closed versus Open Domain Question Answering","year":2007,"lang":"en","type":"article","venue":"International Conference on Semantic Computing (ICSC 2007)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Natural language processing; WordNet; Artificial intelligence; Sentence; Parsing; Feature (linguistics); Semantic similarity; Similarity (geometry); Cosine similarity; Verb; Domain (mathematical analysis); Linguistics; Pattern recognition (psychology); Mathematics; Image (mathematics)","score_opus":0.04951713109689595,"score_gpt":0.3361152564315904,"score_spread":0.2865981253346945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4238310921","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8736098,0.004588832,0.10610775,0.0006965498,0.00019697951,0.00046147488,0.0012567549,0.0027360795,0.010345876],"genre_scores_gemma":[0.9775793,0.00021472239,0.019057414,0.00007949738,0.0001716654,0.000078526864,0.0019339627,0.00015832907,0.000726657],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9800387,0.009994496,0.001515753,0.0023217783,0.0051622414,0.000967071],"domain_scores_gemma":[0.8986353,0.08572519,0.0037155524,0.0030297555,0.007182088,0.0017120095],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014365137,0.0016862729,0.0015634714,0.007677796,0.0007068477,0.0027456146,0.0009560958,0.002028468,0.0015877654],"category_scores_gemma":[0.06889669,0.00034149308,0.0016669339,0.0032414785,0.0011116771,0.004761839,0.0030736919,0.0015955544,0.00093919947],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049369936,0.0012406369,0.18328176,0.0018218774,0.0013915194,0.00083459716,0.002222532,0.026326766,0.051699497,0.002964686,0.003968607,0.71931046],"study_design_scores_gemma":[0.0003355088,0.0054297624,0.4376234,0.0002230312,0.0017658226,0.0031822205,0.0037714199,0.45920357,0.058262054,0.018583722,0.011224664,0.00039476767],"about_ca_topic_score_codex":0.0017203807,"about_ca_topic_score_gemma":0.0014747398,"teacher_disagreement_score":0.014365137,"about_ca_system_score_codex":0.00070766325,"about_ca_system_score_gemma":0.0006970313,"threshold_uncertainty_score":0.07597101},"labels":[],"label_agreement":null},{"id":"W4244315471","doi":"10.4018/978-1-5225-2176-1.ch008","title":"Concept Science","year":2017,"lang":"en","type":"book-chapter","venue":"Advances in computational intelligence and robotics book series","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Sublanguage; Computer science; Meaning (existential); Context (archaeology); Conceptual change; Knowledge management; Sociology; Pedagogy; Epistemology; Artificial intelligence","score_opus":0.034259276769082615,"score_gpt":0.306242618827374,"score_spread":0.2719833420582914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244315471","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008750232,0.028487757,0.019928405,0.0034557446,0.003614633,0.00010504029,0.0003945982,0.0003855669,0.94275326],"genre_scores_gemma":[0.023693478,0.040126614,0.02285425,0.004207109,0.0026318906,0.00023019125,0.0012985084,0.0007227543,0.9042352],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99950814,0.00008110102,0.00002446206,0.00012182556,0.00021913988,0.000045396642],"domain_scores_gemma":[0.9996216,0.00015734757,0.000018182838,0.00005752745,0.00010160277,0.00004359915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00040322045,0.0011041573,0.00047430443,0.002153948,0.0017304075,0.0045266286,0.0009016284,0.0012996522,0.07093254],"category_scores_gemma":[0.0015406932,0.00028608824,0.000684884,0.0019715314,0.0025413625,0.007489142,0.0019144617,0.0030884086,0.026860075],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00000807088,0.000024541157,0.00006622512,0.00028147595,0.000004867482,0.00006603579,0.0009056361,0.0002345134,0.00033795187,0.73653513,0.16047212,0.10106341],"study_design_scores_gemma":[8.8595505e-7,0.0000032641099,0.00003995838,0.0000930108,8.7326407e-7,0.00009904245,0.000085418935,0.000073353476,0.00006798232,0.038980298,0.9605533,0.000002603115],"about_ca_topic_score_codex":0.0017019112,"about_ca_topic_score_gemma":0.0031195574,"teacher_disagreement_score":0.07093254,"about_ca_system_score_codex":0.0027163543,"about_ca_system_score_gemma":0.0016102003,"threshold_uncertainty_score":0.237293},"labels":[],"label_agreement":null},{"id":"W4244795915","doi":"10.32920/ryerson.14654349","title":"Identifying User Interests In An Online Discussion Forum With Deep Learning","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Ontario Tech University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Topic model; Artificial neural network; Metric (unit); Set (abstract data type); Sample (material); Artificial intelligence; Probabilistic logic; Machine learning; Social media; Recommender system; Test set; Data set; Data mining; World Wide Web; Engineering","score_opus":0.05758956287715514,"score_gpt":0.3124403340316915,"score_spread":0.25485077115453636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244795915","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8202277,0.00028752454,0.17261006,0.0004896334,0.00004781389,0.00021759386,0.0008688229,0.001290033,0.003960858],"genre_scores_gemma":[0.9647848,0.000042518124,0.032620087,0.00003559773,0.000012849975,0.00008693241,0.0005020997,0.00002202826,0.0018929844],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987405,0.0005914906,0.000046782523,0.00029446435,0.00016820095,0.00015854013],"domain_scores_gemma":[0.99526274,0.003257358,0.00032928138,0.00030609677,0.00059361954,0.0002510077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003897632,0.0006347803,0.00043422714,0.001758392,0.00061972876,0.0013069386,0.00078715984,0.0009584333,0.0013162888],"category_scores_gemma":[0.00879554,0.0004353188,0.00058346335,0.0009860641,0.00040621168,0.0022195915,0.001420602,0.0016005208,0.00061366806],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025703865,0.002870197,0.15296829,0.0005094622,0.0003772162,0.00053403754,0.00891769,0.26738402,0.038046665,0.01428694,0.0076489584,0.50388616],"study_design_scores_gemma":[0.00001225359,0.00008140787,0.006075306,0.000011155521,0.000012881516,0.000022224156,0.0001931779,0.9864362,0.0028380991,0.0035877344,0.00071511074,0.000014505789],"about_ca_topic_score_codex":0.00680641,"about_ca_topic_score_gemma":0.011204113,"teacher_disagreement_score":0.00680641,"about_ca_system_score_codex":0.0013075715,"about_ca_system_score_gemma":0.0007580225,"threshold_uncertainty_score":0.020612895},"labels":[],"label_agreement":null},{"id":"W4248359551","doi":"10.1007/978-1-4939-7131-2_101307","title":"Tag-Based Recommendation","year":2018,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Information retrieval","score_opus":0.04725276098945988,"score_gpt":0.2523889833887339,"score_spread":0.205136222399274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4248359551","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012083711,0.026138086,0.8369477,0.002139395,0.0032257696,0.00023129719,0.0041238866,0.010359524,0.104750656],"genre_scores_gemma":[0.17246568,0.030987443,0.39292127,0.0010802835,0.0031506056,0.00026684886,0.019206285,0.0020685785,0.3778531],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993637,0.0001075579,0.000032915683,0.000184656,0.00026357206,0.000047712332],"domain_scores_gemma":[0.9991574,0.00024814758,0.000029696359,0.0002701048,0.00025004614,0.000044706132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007239563,0.0014110041,0.0014210234,0.0020792226,0.0005689985,0.0024089438,0.0015769986,0.0011404668,0.02565118],"category_scores_gemma":[0.0021535186,0.00058538036,0.0012230229,0.0048299073,0.00025078573,0.0029313844,0.0008008521,0.001560005,0.03756534],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014345118,0.00022154445,0.0006762226,0.00032541392,0.00017081646,0.0000493809,0.0000461879,0.016068425,0.006487401,0.012901886,0.153854,0.8090552],"study_design_scores_gemma":[0.000071259645,0.00017716788,0.0029050652,0.00022635657,0.00030084662,0.000554809,0.00009585696,0.6229317,0.012639437,0.0639104,0.29602954,0.00015745859],"about_ca_topic_score_codex":0.005435897,"about_ca_topic_score_gemma":0.008834569,"teacher_disagreement_score":0.02565118,"about_ca_system_score_codex":0.0007516107,"about_ca_system_score_gemma":0.0009129798,"threshold_uncertainty_score":0.085811794},"labels":[],"label_agreement":null},{"id":"W4248786004","doi":"10.18653/v1/2021.iwpt-1","title":"Proceedings of the 17th International Conference on Parsing Technologies and the IWPT 2021 Shared Task on Parsing into Enhanced Universal Dependencies (IWPT 2021)","year":2021,"lang":"en","type":"paratext","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Parsing; Computer science; Task (project management); Artificial intelligence; Natural language processing; Engineering","score_opus":0.023962946445698168,"score_gpt":0.25047479970310116,"score_spread":0.226511853257403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4248786004","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013062485,0.03458031,0.6955969,0.04072305,0.036896963,0.00099886,0.013855362,0.03320862,0.13107754],"genre_scores_gemma":[0.053364765,0.025489185,0.36714724,0.006335388,0.008573427,0.0015423687,0.08311844,0.013795648,0.44063342],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962768,0.0013015682,0.00027322918,0.0007793981,0.0009904089,0.00037859773],"domain_scores_gemma":[0.9908896,0.0028499593,0.00020036074,0.0024457842,0.0021703152,0.0014440302],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007273116,0.0020366856,0.0022645278,0.0020446086,0.001583578,0.0069897417,0.0035011473,0.0030368143,0.0912774],"category_scores_gemma":[0.010723972,0.0010402645,0.001564946,0.0031758784,0.0018740989,0.011607068,0.006154876,0.006806706,0.050336495],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033647625,0.00026441284,0.00032631966,0.00034156124,0.000084704625,0.00022166458,0.00027828434,0.0015453233,0.0031694865,0.016052369,0.7862273,0.19115213],"study_design_scores_gemma":[0.0000710943,0.00010042908,0.0007652639,0.0002502985,0.000064581705,0.00029850635,0.00018718246,0.015185143,0.003840475,0.030630426,0.94855857,0.000048043927],"about_ca_topic_score_codex":0.009298765,"about_ca_topic_score_gemma":0.012342138,"teacher_disagreement_score":0.0912774,"about_ca_system_score_codex":0.0024170969,"about_ca_system_score_gemma":0.0039261286,"threshold_uncertainty_score":0.30535328},"labels":[],"label_agreement":null},{"id":"W4248854659","doi":"10.2196/preprints.18055","title":"Exploring the Privacy-Preserving Properties of Word Embeddings: Algorithmic Validation Study (Preprint)","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Michael's Hospital; Institute for Clinical Evaluative Sciences; Vector Institute; University of Toronto","funders":"","keywords":"Word (group theory); Computer science; Preprint; Natural language processing; Word embedding; Code (set theory); Embedding; Information retrieval; Representation (politics); Artificial intelligence; Linguistics; World Wide Web","score_opus":0.20499366790104462,"score_gpt":0.2940062204409732,"score_spread":0.08901255253992857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4248854659","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72915417,0.0036350284,0.24827713,0.0027333503,0.00089493144,0.00075953867,0.0056786574,0.0025809223,0.0062862765],"genre_scores_gemma":[0.8340975,0.0007230957,0.14458606,0.0006214345,0.00017162849,0.0006740148,0.015930379,0.0003056915,0.0028901717],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9893365,0.006910838,0.0008613605,0.0011692955,0.0013204542,0.00040163094],"domain_scores_gemma":[0.9251485,0.056263406,0.002272674,0.010809181,0.0050869575,0.00041936373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014469718,0.0011753764,0.0007156753,0.0013037664,0.0009424687,0.0018934536,0.0014424263,0.0019224837,0.003111186],"category_scores_gemma":[0.07338292,0.00044020134,0.0013863926,0.0013292808,0.0018675907,0.0034838712,0.0032076384,0.002711659,0.0014657932],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027524186,0.0019700734,0.064261794,0.0025096503,0.0012204713,0.0009378638,0.0027198046,0.36967948,0.012371753,0.028989393,0.050595067,0.4619922],"study_design_scores_gemma":[0.00030192407,0.0007596178,0.007944954,0.00028978768,0.00016445693,0.00075825414,0.0009766745,0.9375428,0.014373672,0.028408285,0.008397164,0.000082501],"about_ca_topic_score_codex":0.0029899091,"about_ca_topic_score_gemma":0.0031448922,"teacher_disagreement_score":0.014469718,"about_ca_system_score_codex":0.0012382423,"about_ca_system_score_gemma":0.0015460713,"threshold_uncertainty_score":0.07652408},"labels":[],"label_agreement":null},{"id":"W4251776334","doi":"10.1007/978-0-387-39940-9_3801","title":"Text/Document Summarization","year":2009,"lang":"en","type":"book-chapter","venue":"Encyclopedia of Database Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Multi-document summarization; Information retrieval; Natural language processing","score_opus":0.015605505967897022,"score_gpt":0.23223643188346607,"score_spread":0.21663092591556904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4251776334","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011795441,0.013260605,0.80412924,0.002142381,0.0027739503,0.0014984383,0.019976458,0.0602678,0.084155634],"genre_scores_gemma":[0.05428875,0.010596982,0.6532791,0.000925261,0.0022860193,0.0008236202,0.06835935,0.0064918874,0.20294903],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993742,0.00008770712,0.000084432024,0.00018686296,0.00021687134,0.00004994028],"domain_scores_gemma":[0.99840766,0.00033301234,0.00011689454,0.0002986715,0.00078258914,0.00006121142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00081221457,0.001825718,0.0013087995,0.0043545524,0.00077899185,0.0025540749,0.0014005624,0.00068427273,0.040577944],"category_scores_gemma":[0.0023914361,0.0004303989,0.0009698314,0.004179177,0.00027681058,0.0021830238,0.00138107,0.0009901229,0.043830596],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008672601,0.00006427062,0.00018051497,0.0007927252,0.000055421016,0.000117281765,0.000121833786,0.0010422054,0.024797412,0.0025170972,0.11219542,0.85802907],"study_design_scores_gemma":[0.00007122616,0.00032091062,0.0027648518,0.00036057492,0.00050160603,0.0011446093,0.00052585156,0.033105012,0.13178262,0.015281312,0.8140221,0.00011931175],"about_ca_topic_score_codex":0.0009028255,"about_ca_topic_score_gemma":0.0015587927,"teacher_disagreement_score":0.040577944,"about_ca_system_score_codex":0.00035482735,"about_ca_system_score_gemma":0.00082232954,"threshold_uncertainty_score":0.13574678},"labels":[],"label_agreement":null},{"id":"W4252349324","doi":"10.3115/1572306.1572315","title":"Recognizing speculative language in biomedical research articles","year":2008,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Speculation; Computer science; Weighting; Point (geometry); Scheme (mathematics); Simple (philosophy); Natural language processing; Artificial intelligence; Work (physics); Linguistics; Epistemology; Mathematics; Finance","score_opus":0.1893865474279963,"score_gpt":0.3802550046572641,"score_spread":0.1908684572292678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252349324","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41011056,0.0043282523,0.5703751,0.004377099,0.00041274406,0.0003339045,0.0017090888,0.0015076565,0.006845645],"genre_scores_gemma":[0.7736235,0.0013433319,0.2201635,0.00040049196,0.00075340655,0.00022568917,0.0016799802,0.00011865421,0.0016915072],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9981098,0.00069045904,0.00029132218,0.0003386542,0.00047656623,0.00009324092],"domain_scores_gemma":[0.9677173,0.022929104,0.0052382764,0.0012596045,0.0024721208,0.00038365243],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.003974359,0.00061522704,0.00056963204,0.0064691263,0.0008893171,0.0024065867,0.000962971,0.0010174958,0.001114975],"category_scores_gemma":[0.021327375,0.00041504705,0.00065843976,0.004267568,0.0011373492,0.004377606,0.001472233,0.0010782203,0.000408762],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00095212436,0.00028703897,0.07064231,0.0028079392,0.00039826604,0.005505123,0.015566471,0.02283982,0.11544919,0.13477907,0.012881842,0.61789083],"study_design_scores_gemma":[0.000113294715,0.00065133086,0.05229104,0.00068124867,0.00065388635,0.0054541826,0.0056568678,0.4954907,0.06938306,0.30336076,0.065935284,0.00032832785],"about_ca_topic_score_codex":0.00088279304,"about_ca_topic_score_gemma":0.0013575862,"teacher_disagreement_score":0.9960256,"about_ca_system_score_codex":0.0005654541,"about_ca_system_score_gemma":0.0010519114,"threshold_uncertainty_score":0.021018684},"labels":[],"label_agreement":null},{"id":"W4254133819","doi":"10.1145/564437.564448","title":"The impact of corpus size on question answering performance","year":2002,"lang":"en","type":"article","venue":"Proceedings of the 25th annual international ACM SIGIR conference on Research and development in information retrieval - SIGIR '02","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Terabyte; Question answering; Computer science; Information retrieval; Natural language processing; World Wide Web","score_opus":0.06702928582532092,"score_gpt":0.3316985989911562,"score_spread":0.2646693131658353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4254133819","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9261341,0.0102453185,0.023233505,0.004227948,0.0013078874,0.00058971625,0.0059996024,0.013847632,0.014414242],"genre_scores_gemma":[0.93317044,0.0030646466,0.033708308,0.0010185569,0.0007832691,0.00093736727,0.018766956,0.0027144465,0.005836045],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9827697,0.00837652,0.0023412104,0.0027797644,0.0030726432,0.0006601609],"domain_scores_gemma":[0.8419714,0.13283053,0.0019770518,0.00995807,0.011291098,0.001971887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017559476,0.0013528592,0.0022314824,0.0024405664,0.0021149304,0.0048062936,0.0020659114,0.0021194143,0.0054000164],"category_scores_gemma":[0.1470115,0.0013224347,0.0008904939,0.0036473158,0.0016847694,0.009566499,0.0032884127,0.0019250554,0.0033701085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01643308,0.0033063067,0.052166853,0.005376299,0.0013010096,0.0019165907,0.005151569,0.060516972,0.15777299,0.0045145117,0.10228458,0.58925927],"study_design_scores_gemma":[0.0035107185,0.012354464,0.08523628,0.0008005277,0.003040988,0.0052250046,0.0056125014,0.49392703,0.29175198,0.01756523,0.08010968,0.0008655268],"about_ca_topic_score_codex":0.006428584,"about_ca_topic_score_gemma":0.006373196,"teacher_disagreement_score":0.017559476,"about_ca_system_score_codex":0.001198404,"about_ca_system_score_gemma":0.0019921265,"threshold_uncertainty_score":0.09286451},"labels":[],"label_agreement":null},{"id":"W4255223161","doi":"10.18653/v1/w18-26","title":"Proceedings of the Workshop on Machine Reading for Question Answering","year":2018,"lang":"en","type":"paratext","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"Samsung; Tencent","keywords":"Question answering; Computer science; Reading (process); Artificial intelligence; Information retrieval; Natural language processing; World Wide Web; Linguistics; Philosophy","score_opus":0.027796324734331357,"score_gpt":0.28964389893517956,"score_spread":0.2618475742008482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255223161","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025884615,0.062304847,0.766734,0.053536825,0.015229849,0.0013144343,0.00798717,0.016946016,0.050062332],"genre_scores_gemma":[0.18044847,0.023171136,0.59015054,0.007789897,0.009667899,0.0025569946,0.04520886,0.0043228986,0.13668332],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99442935,0.0031949708,0.00040162305,0.0011165127,0.0005962883,0.000261295],"domain_scores_gemma":[0.98630583,0.008657897,0.00020677807,0.0026136234,0.001512367,0.00070350344],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012890696,0.0024897275,0.0024813288,0.0021141698,0.0011716338,0.007873029,0.004925936,0.0036734443,0.050630126],"category_scores_gemma":[0.019822849,0.0010484722,0.002472495,0.0017375653,0.0018073795,0.011833372,0.0043331226,0.005952725,0.020180788],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063417025,0.0005399715,0.0008442191,0.00096149597,0.00031671414,0.00036132222,0.0010536961,0.0056356816,0.00523644,0.02707508,0.41938472,0.53795654],"study_design_scores_gemma":[0.00020180238,0.00027444624,0.003175212,0.00052543415,0.00021125506,0.0005740031,0.00077528035,0.13098851,0.006625983,0.09892728,0.7576163,0.000104537365],"about_ca_topic_score_codex":0.0053788326,"about_ca_topic_score_gemma":0.0058580446,"teacher_disagreement_score":0.050630126,"about_ca_system_score_codex":0.0027573667,"about_ca_system_score_gemma":0.002436633,"threshold_uncertainty_score":0.16937464},"labels":[],"label_agreement":null},{"id":"W4255436026","doi":"10.24124/2018/58808","title":"Sentence encoders for semantic textual similarity: A survey.","year":2018,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Sentence; Semantic similarity; Semantics (computer science); Task (project management); Similarity (geometry); Representation (politics); Focus (optics); Semantic role labeling; Encoder; Word (group theory); Distributional semantics; Deep learning; Linguistics; Programming language","score_opus":0.04599552660613468,"score_gpt":0.3053835383355127,"score_spread":0.259388011729378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255436026","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043671153,0.10175947,0.7795475,0.0027684772,0.001870358,0.0005779488,0.012159363,0.025163848,0.032481886],"genre_scores_gemma":[0.41604292,0.05784184,0.44642782,0.001853538,0.0016906159,0.0009090925,0.04644471,0.00143186,0.027357645],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99939907,0.0001627692,0.00006559858,0.00014584653,0.00019043207,0.000036228877],"domain_scores_gemma":[0.99861336,0.00066197006,0.000105410916,0.00022423115,0.00033451282,0.000060539613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012905205,0.0009834152,0.000849446,0.0023042941,0.00026709138,0.00091740343,0.0015866357,0.0006929414,0.009351083],"category_scores_gemma":[0.00531526,0.0003736987,0.0007103708,0.0017413659,0.00026060402,0.0032836806,0.0009192812,0.0009163137,0.0053143194],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011554633,0.00017412318,0.0011816574,0.0009065501,0.00010155015,0.00009159447,0.000085733045,0.005033813,0.0035541528,0.010815092,0.040595297,0.9373449],"study_design_scores_gemma":[0.00010846114,0.0007770638,0.008666161,0.0012623869,0.0004600102,0.0017901757,0.0004951793,0.690491,0.030018963,0.07575964,0.19004185,0.00012915308],"about_ca_topic_score_codex":0.0033117535,"about_ca_topic_score_gemma":0.0040941546,"teacher_disagreement_score":0.009351083,"about_ca_system_score_codex":0.0006983515,"about_ca_system_score_gemma":0.0010578552,"threshold_uncertainty_score":0.031282485},"labels":[],"label_agreement":null},{"id":"W4255582338","doi":"10.24124/2017/1381","title":"Architecture for automatic poetry generation through pattern recognition","year":2017,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"","keywords":"Computer science; Artificial intelligence; Architecture; Cluster analysis; Representation (politics); Natural language processing; Inference; Unsupervised learning; Agile software development; Stability (learning theory); Computational linguistics; Machine learning; Software engineering","score_opus":0.0610981320155035,"score_gpt":0.3110912567792143,"score_spread":0.24999312476371077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255582338","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016600696,0.0002864297,0.9514587,0.00026944716,0.00016575241,0.00014790063,0.00040463122,0.02465856,0.0060079484],"genre_scores_gemma":[0.21279474,0.0006404022,0.7608364,0.00023071679,0.0001378159,0.00039751516,0.0037398746,0.00090037537,0.020322204],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998024,0.000021023794,0.000015188848,0.000081813894,0.000052827552,0.000026796479],"domain_scores_gemma":[0.99969923,0.000059932856,0.000016245971,0.00008357172,0.00011697635,0.000024049083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000457054,0.000570695,0.0005125011,0.00088020443,0.0005690495,0.001384986,0.0018179016,0.00068839284,0.007559358],"category_scores_gemma":[0.0010159498,0.00040498519,0.0008560691,0.0009033633,0.00036226565,0.0020239817,0.00093445013,0.0010100319,0.0071249255],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001730099,0.00027991156,0.0016477309,0.00024824982,0.00011074851,0.00018277252,0.00037304978,0.019969963,0.060382362,0.017632512,0.027006865,0.87199277],"study_design_scores_gemma":[0.00004049125,0.00016769029,0.0018614638,0.000042347463,0.00011183149,0.0003527312,0.00016511143,0.84537953,0.06536507,0.029624797,0.056847584,0.000041303334],"about_ca_topic_score_codex":0.0025099996,"about_ca_topic_score_gemma":0.0038810223,"teacher_disagreement_score":0.007559358,"about_ca_system_score_codex":0.00054665044,"about_ca_system_score_gemma":0.00081286137,"threshold_uncertainty_score":0.025288582},"labels":[],"label_agreement":null},{"id":"W4256072618","doi":"10.31234/osf.io/ytnjp","title":"Indirect associations in learning semantic and syntactic lexical relationships","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Carleton University","funders":"","keywords":"Natural language processing; Artificial intelligence; Computer science; Construct (python library); Word (group theory); Associative property; Semantics (computer science); Association (psychology); Word Association; Part of speech; Similarity (geometry); Meaning (existential); Linguistics; Psychology; Mathematics","score_opus":0.06851572527720474,"score_gpt":0.2848876306803419,"score_spread":0.21637190540313717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256072618","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70109713,0.00035717775,0.29441583,0.00047167047,0.000026992655,0.00003433494,0.000104873005,0.0003110163,0.003181043],"genre_scores_gemma":[0.9701488,0.000194834,0.027921334,0.000054304637,0.000021091111,0.0000658448,0.00020370814,0.000033099688,0.001357016],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993173,0.00032120026,0.000047995978,0.0001844475,0.0000784382,0.00005070507],"domain_scores_gemma":[0.994572,0.004149004,0.00041344983,0.00046247232,0.00024521464,0.00015786986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015408187,0.0005006946,0.000671313,0.0008099164,0.0004872562,0.0012815185,0.0009396685,0.000718104,0.0016220447],"category_scores_gemma":[0.0125386845,0.00051060267,0.000672249,0.0009074973,0.0011056525,0.0050562406,0.0021500958,0.0013620626,0.00029055946],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069160014,0.0004030905,0.037840623,0.00037558962,0.00044462332,0.0002926782,0.0016286452,0.4958302,0.008934199,0.165841,0.0015779925,0.2861398],"study_design_scores_gemma":[0.000028640008,0.00010391244,0.002537598,0.000017232673,0.000050206432,0.00008779494,0.00009979514,0.7402333,0.0023337596,0.25388968,0.0005985037,0.000019541038],"about_ca_topic_score_codex":0.0014273551,"about_ca_topic_score_gemma":0.0019119614,"teacher_disagreement_score":0.0016220447,"about_ca_system_score_codex":0.0006489992,"about_ca_system_score_gemma":0.0004946136,"threshold_uncertainty_score":0.00814873},"labels":[],"label_agreement":null},{"id":"W4280612849","doi":"10.18653/v1/2022.acl-short.45","title":"Problems with Cosine as a Measure of Embedding Similarity for High Frequency Words","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Cosine similarity; Similarity (geometry); Polysemy; Word (group theory); Natural language processing; Trigonometric functions; Embedding; Computer science; Argument (complex analysis); Artificial intelligence; Mathematics; Conjecture; Discrete cosine transform; TRACE (psycholinguistics); Speech recognition; Pattern recognition (psychology); Linguistics; Pure mathematics","score_opus":0.02659254907497569,"score_gpt":0.2515398599478576,"score_spread":0.22494731087288192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4280612849","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05839867,0.0064801024,0.92264,0.0042068386,0.0007548587,0.000275369,0.0011939101,0.0007269094,0.005323303],"genre_scores_gemma":[0.6896952,0.0028288364,0.29761285,0.0012496001,0.0013964926,0.00095678505,0.0025966105,0.0004971903,0.0031665291],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9463965,0.02757302,0.0060680225,0.009639301,0.00951334,0.0008098026],"domain_scores_gemma":[0.838033,0.1232817,0.008644593,0.02039782,0.008679267,0.0009637548],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032326654,0.0014434871,0.0031463993,0.005434036,0.0023595686,0.006237093,0.0041037276,0.0031783232,0.0020325098],"category_scores_gemma":[0.18755938,0.0012577142,0.0013315609,0.0114098275,0.0047384286,0.010283471,0.0042972676,0.0050458284,0.0017550881],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078470446,0.0003240744,0.042451944,0.0027604757,0.0013248249,0.00053650036,0.0039395755,0.06905963,0.009293888,0.29577178,0.03474575,0.53900695],"study_design_scores_gemma":[0.00007775119,0.00030120683,0.020444605,0.00048747624,0.00014387895,0.0017113781,0.0020640776,0.31689432,0.007764201,0.62677455,0.023016507,0.00032008905],"about_ca_topic_score_codex":0.005865411,"about_ca_topic_score_gemma":0.0030486407,"teacher_disagreement_score":0.032326654,"about_ca_system_score_codex":0.0026975446,"about_ca_system_score_gemma":0.0013966843,"threshold_uncertainty_score":0.17096174},"labels":[],"label_agreement":null},{"id":"W4281263937","doi":"10.18653/v1/2022.findings-acl.117","title":"Two-Step Question Retrieval for Open-Domain QA","year":2022,"lang":"en","type":"article","venue":"Findings of the Association for Computational Linguistics: ACL 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Ministry of Science and ICT, South Korea; National Research Foundation of Korea; National Research Foundation","keywords":"Search engine indexing; Inference; Computer science; Information retrieval; Labrador Retriever; Pipeline (software); Squid; Domain (mathematical analysis); Artificial intelligence; Open domain; Question answering; Mathematics; Programming language","score_opus":0.01875750815666395,"score_gpt":0.2925474942045004,"score_spread":0.27378998604783644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281263937","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027755762,0.0010419132,0.95185965,0.0008017536,0.0001089852,0.00037247493,0.0006811127,0.014229616,0.0031487243],"genre_scores_gemma":[0.48502052,0.00048530745,0.5036844,0.0005585844,0.00016379434,0.0003594253,0.0033241517,0.00045971698,0.005944135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985671,0.00059229607,0.00008777731,0.00040124197,0.0002341387,0.000117505864],"domain_scores_gemma":[0.99486226,0.0026822905,0.00016620637,0.0013664119,0.00071025547,0.00021262499],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034847767,0.000825689,0.0010310212,0.0013334843,0.00067610055,0.0011986323,0.002568392,0.0017909621,0.009398755],"category_scores_gemma":[0.010890025,0.00066473644,0.0012298697,0.0009476984,0.0009235206,0.005374018,0.002892502,0.0028009457,0.005467695],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001002704,0.0008574103,0.0052059474,0.00095487264,0.00018287523,0.0002082218,0.00086227385,0.07719907,0.025425572,0.028891327,0.03443011,0.82477957],"study_design_scores_gemma":[0.00008736369,0.00019659585,0.00081500714,0.000021865451,0.000036723184,0.00019679489,0.00007878431,0.96172875,0.010281356,0.020159071,0.006366363,0.000031284897],"about_ca_topic_score_codex":0.008352006,"about_ca_topic_score_gemma":0.010537689,"teacher_disagreement_score":0.009398755,"about_ca_system_score_codex":0.0012339873,"about_ca_system_score_gemma":0.0021475325,"threshold_uncertainty_score":0.031441987},"labels":[],"label_agreement":null},{"id":"W4281482237","doi":"10.1145/3534965","title":"How to Approach Ambiguous Queries in Conversational Search: A Survey of Techniques, Approaches, Tools, and Challenges","year":2022,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Focus (optics); Conversation; Context (archaeology); Process (computing); Task (project management); World Wide Web; Natural language; Human–computer interaction; Information retrieval; Artificial intelligence","score_opus":0.4837335270612913,"score_gpt":0.33671637581225367,"score_spread":0.14701715124903764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281482237","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002561632,0.91653866,0.06819654,0.0041240514,0.00042324932,0.00015480304,0.00015701125,0.00040934372,0.0074346997],"genre_scores_gemma":[0.03728814,0.8430406,0.111252636,0.0017727733,0.0013730571,0.00025305024,0.0007466706,0.00020508685,0.004067974],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99709475,0.001112167,0.00032174727,0.00047000695,0.00081524695,0.0001860797],"domain_scores_gemma":[0.98971176,0.007298253,0.00036146757,0.00064708263,0.0018155375,0.00016587343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003902745,0.0012439111,0.0017313897,0.0055850153,0.0009846236,0.003992049,0.0020922928,0.0019493771,0.0033248763],"category_scores_gemma":[0.011029355,0.00081122445,0.0013472275,0.0065045157,0.0014442214,0.010640738,0.0018597629,0.002103143,0.00405658],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047984948,0.0001102533,0.0011617343,0.00827523,0.000080611666,0.00008569748,0.00094919646,0.00148498,0.0015942368,0.017766634,0.018784989,0.9496584],"study_design_scores_gemma":[0.000050742776,0.00022311008,0.00350853,0.008622879,0.0004055984,0.0021921757,0.005246331,0.027145574,0.003570023,0.0824779,0.86632574,0.0002314495],"about_ca_topic_score_codex":0.0057095443,"about_ca_topic_score_gemma":0.005129228,"teacher_disagreement_score":0.0057095443,"about_ca_system_score_codex":0.0010427802,"about_ca_system_score_gemma":0.002998152,"threshold_uncertainty_score":0.020639956},"labels":[],"label_agreement":null},{"id":"W4281626616","doi":"10.3390/info13060290","title":"Contextualizer: Connecting the Dots of Context with Second-Order Attention","year":2022,"lang":"en","type":"article","venue":"Information","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Government of Canada; Compute Canada","keywords":"Security token; Computer science; Sentence; Representation (politics); Theoretical computer science; Embedding; Computation; Transformer; Algorithm; Artificial intelligence","score_opus":0.0133513662740684,"score_gpt":0.22015760679256827,"score_spread":0.20680624051849986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281626616","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024362229,0.00038592107,0.96882844,0.0003846468,0.00009275497,0.00007126592,0.00009370882,0.002166467,0.0036145994],"genre_scores_gemma":[0.64218,0.0004857417,0.35080716,0.00030758767,0.00013979852,0.00015895226,0.00023563793,0.00036353566,0.005321584],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99945515,0.00019458303,0.000018808363,0.00020720626,0.000067236666,0.0000570486],"domain_scores_gemma":[0.99899465,0.00051169295,0.000085352585,0.00024485064,0.000099479614,0.000064006716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00095617224,0.00088192546,0.00059292157,0.0009540079,0.00052992883,0.0014609165,0.0014422127,0.0008971803,0.005895558],"category_scores_gemma":[0.0041192397,0.000631109,0.000855209,0.00089901075,0.0010148813,0.004326313,0.0020761865,0.0013625531,0.0009277292],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007668785,0.00022684001,0.0038170442,0.00034267866,0.0001968812,0.00031850487,0.0013035424,0.08013682,0.031632762,0.19848676,0.0075901,0.6751812],"study_design_scores_gemma":[0.00004303419,0.00020557796,0.0015120619,0.000038524275,0.00011110827,0.00019110672,0.000103940874,0.80534065,0.013693973,0.17094155,0.007777225,0.00004122828],"about_ca_topic_score_codex":0.0047550467,"about_ca_topic_score_gemma":0.006913493,"teacher_disagreement_score":0.005895558,"about_ca_system_score_codex":0.00085797015,"about_ca_system_score_gemma":0.00078942504,"threshold_uncertainty_score":0.01972258},"labels":[],"label_agreement":null},{"id":"W4281652650","doi":"10.2196/38052","title":"Exploiting Intersentence Information for Better Question-Driven Abstractive Summarization: Algorithm Development and Validation","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Natural Science Foundation of Liaoning Province; National Natural Science Foundation of China","keywords":"Computer science; Automatic summarization; Development (topology); Artificial intelligence; Algorithm; Natural language processing; Mathematics","score_opus":0.020041629270150883,"score_gpt":0.27434353950104834,"score_spread":0.25430191023089743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281652650","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024124261,0.00093873334,0.96524626,0.00039594262,0.0000809673,0.00059403636,0.00042334874,0.0075361403,0.0006602919],"genre_scores_gemma":[0.20367526,0.00076196116,0.7876656,0.00035383157,0.00012164144,0.0008803983,0.0046480065,0.0003418801,0.0015514757],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981421,0.0006427766,0.00019516949,0.00058218493,0.0003013633,0.00013632135],"domain_scores_gemma":[0.99428236,0.0030742206,0.0002868924,0.0005274208,0.0016529219,0.00017628467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004815519,0.0017067232,0.0013542753,0.0021900346,0.00060622965,0.0016584079,0.0020599999,0.0018075997,0.003760294],"category_scores_gemma":[0.012104986,0.00048219107,0.0013170907,0.0015850494,0.00053635816,0.0030683284,0.0016016557,0.0022856023,0.002332417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003928722,0.0004407729,0.0035943994,0.000538259,0.00027894008,0.0001635383,0.0003237283,0.12186438,0.02306143,0.0030267574,0.008478888,0.83783615],"study_design_scores_gemma":[0.00006255912,0.000121607925,0.00064675836,0.000022534561,0.0000683579,0.000058212998,0.000073955736,0.9877543,0.007430441,0.002049471,0.0016936887,0.000018166702],"about_ca_topic_score_codex":0.00683988,"about_ca_topic_score_gemma":0.0067559015,"teacher_disagreement_score":0.00683988,"about_ca_system_score_codex":0.001205627,"about_ca_system_score_gemma":0.0027006932,"threshold_uncertainty_score":0.025467217},"labels":[],"label_agreement":null},{"id":"W4281741206","doi":"10.1145/3532106.3533506","title":"From Tool to Companion: Storywriters Want AI Writers to Respect Their Personal Values and Writing Strategies","year":2022,"lang":"en","type":"article","venue":"Designing Interactive Systems Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":105,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; World Wide Web","score_opus":0.05415452581127535,"score_gpt":0.292711023264157,"score_spread":0.23855649745288166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281741206","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69710845,0.0007576531,0.16549908,0.0057277526,0.00037297298,0.0003986228,0.00022087454,0.0022746779,0.12763987],"genre_scores_gemma":[0.9557039,0.00014817396,0.028017709,0.0008295047,0.00007371568,0.00021812445,0.00011031641,0.00030184913,0.014596694],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.99336445,0.0042989952,0.0003057269,0.0006699587,0.0010284116,0.00033239965],"domain_scores_gemma":[0.97321486,0.014621015,0.0025395341,0.0041550905,0.002520913,0.0029486138],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.00674058,0.0008163308,0.00044135324,0.0008972094,0.0034849208,0.008604135,0.0009782877,0.0015919147,0.0072294637],"category_scores_gemma":[0.036793664,0.00038243225,0.00036550045,0.00048609715,0.0044815186,0.007065779,0.0053168456,0.0018923646,0.0029923706],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005464874,0.00042433952,0.028909568,0.000871115,0.000065706976,0.0025290246,0.74571604,0.00058937253,0.03217788,0.029592535,0.01399377,0.1445842],"study_design_scores_gemma":[0.00013907115,0.0012610374,0.025440512,0.0007728206,0.00013660884,0.010009947,0.36708394,0.0067044483,0.029148791,0.039477482,0.5195348,0.00029053944],"about_ca_topic_score_codex":0.00041673836,"about_ca_topic_score_gemma":0.00047462483,"teacher_disagreement_score":0.9965151,"about_ca_system_score_codex":0.0008090378,"about_ca_system_score_gemma":0.00091066887,"threshold_uncertainty_score":0.035648048},"labels":[],"label_agreement":null},{"id":"W4281742449","doi":"10.1016/j.health.2022.100068","title":"A COVID-19 Search Engine (CO-SE) with Transformer-based architecture","year":2022,"lang":"en","type":"article","venue":"Healthcare Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"","keywords":"Generalizability theory; Computer science; Coronavirus disease 2019 (COVID-19); Information retrieval; Transformer; Benchmark (surveying); Question answering; Natural language processing; Artificial intelligence; Infectious disease (medical specialty); Disease","score_opus":0.0826462659495686,"score_gpt":0.3513920236641443,"score_spread":0.2687457577145757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281742449","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.099184476,0.0059281085,0.74932545,0.0018590649,0.00041969115,0.0018761758,0.012186622,0.11153754,0.017682822],"genre_scores_gemma":[0.4156644,0.002276156,0.52326816,0.0009985355,0.00020062855,0.00070968684,0.036910884,0.0010347107,0.018936882],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993104,0.000094835006,0.00007726623,0.0002321634,0.00019741445,0.000087872315],"domain_scores_gemma":[0.99936,0.00019214506,0.00003607752,0.00011143785,0.00023423122,0.00006606486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011164342,0.0010628063,0.0012908574,0.0039042337,0.00050856016,0.001290989,0.0019902997,0.0014256183,0.004504025],"category_scores_gemma":[0.0026309632,0.00045453664,0.0012166892,0.0030795375,0.0003578267,0.0029476085,0.0014307867,0.00068338186,0.00451003],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018529693,0.0012399154,0.015590757,0.0019619004,0.0005931289,0.0007971632,0.00047433746,0.059851717,0.04268197,0.02046091,0.10324873,0.7512465],"study_design_scores_gemma":[0.00017074482,0.0004712694,0.0026369432,0.000051593845,0.00018983673,0.0007598294,0.00012785051,0.9360723,0.019480506,0.0077540153,0.03219467,0.00009044015],"about_ca_topic_score_codex":0.016895356,"about_ca_topic_score_gemma":0.022033533,"teacher_disagreement_score":0.016895356,"about_ca_system_score_codex":0.0010494225,"about_ca_system_score_gemma":0.0019604408,"threshold_uncertainty_score":0.033594012},"labels":[],"label_agreement":null},{"id":"W4281746134","doi":"10.3233/shti220155","title":"Towards Automated Screening of Literature on Artificial Intelligence in Nursing","year":2022,"lang":"en","type":"review","venue":"Studies in health technology and informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Academy of Finland","keywords":"Computer science; Relevance (law); Artificial intelligence; Set (abstract data type); Focus (optics); Process (computing); Machine learning; Natural language processing; Data science; Information retrieval","score_opus":0.1467779942212781,"score_gpt":0.4490636340029064,"score_spread":0.30228563978162826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281746134","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007205703,0.9750289,0.01285455,0.0010812853,0.00030638697,0.0004430083,0.0002658534,0.00019158784,0.002622833],"genre_scores_gemma":[0.06295653,0.8614067,0.070598155,0.0009089043,0.0005503008,0.0009345354,0.0013867915,0.000047927806,0.0012102098],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9933547,0.0034050387,0.0007645324,0.00034802165,0.00200722,0.000120647695],"domain_scores_gemma":[0.97165424,0.016673991,0.0031574336,0.0006423524,0.007615393,0.00025661712],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.016405603,0.00087806454,0.0017500845,0.018673783,0.00041340775,0.0025678435,0.0011868898,0.00082336017,0.0014913653],"category_scores_gemma":[0.02687251,0.0004853804,0.0014708737,0.009272923,0.00033156143,0.0022216812,0.00097059447,0.00071309163,0.0009805948],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000681772,0.00007676005,0.0019479753,0.0274231,0.0004920447,0.00005243494,0.00015617635,0.0006574411,0.0014660989,0.00066598936,0.004087571,0.96290636],"study_design_scores_gemma":[0.00040969902,0.002502104,0.09608862,0.15234485,0.016540125,0.0026832218,0.0013909888,0.03163573,0.017856453,0.013320225,0.66487193,0.0003560498],"about_ca_topic_score_codex":0.0026094152,"about_ca_topic_score_gemma":0.0054566218,"teacher_disagreement_score":0.9835944,"about_ca_system_score_codex":0.0010883709,"about_ca_system_score_gemma":0.004103494,"threshold_uncertainty_score":0.08676213},"labels":[],"label_agreement":null},{"id":"W4281817583","doi":"10.1186/s12859-022-04751-6","title":"CoQUAD: a COVID-19 question answering dataset system, facilitating research, benchmarking, and practice","year":2022,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Public Health Ontario","funders":"Institute of Health Services and Policy Research; Canadian Institutes of Health Research","keywords":"Benchmarking; Coronavirus disease 2019 (COVID-19); 2019-20 coronavirus outbreak; Data science; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Computer science; Computational biology; Information retrieval; Biology; Virology; Medicine; Business","score_opus":0.12784802972864218,"score_gpt":0.3884872227123956,"score_spread":0.26063919298375343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281817583","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035070743,0.0057313326,0.068453096,0.004244152,0.0011104889,0.0052360995,0.7471319,0.11501745,0.018004721],"genre_scores_gemma":[0.023559693,0.0004982368,0.09231154,0.0011611052,0.00016240362,0.0027087517,0.8754614,0.0011567347,0.0029802267],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99193066,0.0025323178,0.0013470384,0.0018341609,0.0019516649,0.00040413986],"domain_scores_gemma":[0.9882428,0.0042016394,0.0011158742,0.002614662,0.0028881163,0.00093686406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0086199185,0.0029120771,0.0012492738,0.008067946,0.0016976738,0.0028152259,0.0047256546,0.0037536712,0.013236554],"category_scores_gemma":[0.025953213,0.00059207814,0.0017792865,0.004995091,0.0009141659,0.0055985195,0.0060588624,0.0028141527,0.011804751],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012139123,0.00081483094,0.008541645,0.0042738267,0.00026426345,0.0004569204,0.0009345002,0.0036369823,0.009874526,0.006015397,0.8282717,0.13570167],"study_design_scores_gemma":[0.0012409724,0.001012741,0.024866536,0.00083878747,0.00020847496,0.00088273175,0.0015207351,0.08225365,0.02562059,0.014067866,0.8470841,0.0004028225],"about_ca_topic_score_codex":0.015389197,"about_ca_topic_score_gemma":0.022310983,"teacher_disagreement_score":0.015389197,"about_ca_system_score_codex":0.0030436693,"about_ca_system_score_gemma":0.004086695,"threshold_uncertainty_score":0.045587063},"labels":[],"label_agreement":null},{"id":"W4283720480","doi":"10.1007/s10115-022-01703-7","title":"BertHANK: hierarchical attention networks with enhanced knowledge and pre-trained model for answer selection","year":2022,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Selection (genetic algorithm); Computer science; Artificial intelligence; Machine learning; Model selection","score_opus":0.011614359370299665,"score_gpt":0.2366141620126121,"score_spread":0.22499980264231242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283720480","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021390704,0.0017069129,0.92851824,0.0010276764,0.00045552463,0.00026761153,0.0023749839,0.039043337,0.0052149906],"genre_scores_gemma":[0.35670865,0.0007468638,0.6031391,0.0009444948,0.00032142006,0.0006452149,0.0072141383,0.0024442824,0.027835814],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921083,0.00021709803,0.0000310509,0.00026633116,0.00015580465,0.000118830874],"domain_scores_gemma":[0.9985726,0.00084196276,0.000047173355,0.0002039564,0.00024931776,0.00008511741],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016127425,0.0015780132,0.0011772801,0.0013431228,0.0008725575,0.0013050479,0.0031351934,0.0030984713,0.013434308],"category_scores_gemma":[0.0061694155,0.0010293905,0.0010421928,0.0013672658,0.00055154366,0.004161507,0.0029644084,0.003082603,0.0052142455],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010751999,0.0003655997,0.00093218003,0.00035065116,0.0002677906,0.0002190921,0.00025868468,0.09616026,0.0131795015,0.013227494,0.07606995,0.7978936],"study_design_scores_gemma":[0.00006224343,0.00003503421,0.00022232335,0.000013477849,0.000039028135,0.000024828409,0.000018972856,0.97988635,0.0040190043,0.012904226,0.0027554925,0.000019029449],"about_ca_topic_score_codex":0.025516426,"about_ca_topic_score_gemma":0.03904148,"teacher_disagreement_score":0.025516426,"about_ca_system_score_codex":0.0014902877,"about_ca_system_score_gemma":0.0015698806,"threshold_uncertainty_score":0.05073577},"labels":[],"label_agreement":null},{"id":"W4283789759","doi":"10.1609/aaai.v36i10.21332","title":"Search and Learn: Improving Semantic Coverage for Data-to-Text Generation","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Bundesministerium für Bildung und Forschung; Compute Canada; Technische Universität Kaiserslautern; Canadian Institute for Advanced Research; DeepMind; Nvidia","keywords":"Computer science; Inference; Focus (optics); Artificial intelligence; Limiting; Cover (algebra); Quality (philosophy); Training set; Language model; Natural language processing; Information retrieval; Machine learning","score_opus":0.2087035607689111,"score_gpt":0.3334275307102393,"score_spread":0.12472396994132817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283789759","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0775306,0.0018903511,0.8806764,0.00089412014,0.00016628024,0.00028575628,0.0017044714,0.033118717,0.0037334159],"genre_scores_gemma":[0.59798145,0.0005174453,0.38663185,0.00091215665,0.00017625754,0.00043640804,0.0075165136,0.0022464863,0.0035813279],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981027,0.0006669054,0.0001281386,0.00053873076,0.0004309215,0.00013257534],"domain_scores_gemma":[0.99334455,0.0047368235,0.00027162858,0.0009204653,0.0005624467,0.00016399281],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023187094,0.0012786083,0.0015767976,0.0019920385,0.000741813,0.001314774,0.0021720468,0.0017520755,0.004044173],"category_scores_gemma":[0.016027184,0.0006412526,0.0010832177,0.0014287205,0.0010283426,0.0052339407,0.002729187,0.001763876,0.0018334647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011914439,0.0008376741,0.005690843,0.0007875567,0.00018649192,0.00073615764,0.0010920719,0.17206366,0.02135267,0.013670403,0.03878034,0.74361074],"study_design_scores_gemma":[0.000094177754,0.00008900845,0.00037010043,0.000027605,0.00003944569,0.0001416999,0.000083738276,0.9731986,0.009907241,0.012395594,0.0036321438,0.000020684696],"about_ca_topic_score_codex":0.004297231,"about_ca_topic_score_gemma":0.006493174,"teacher_disagreement_score":0.004297231,"about_ca_system_score_codex":0.0010167466,"about_ca_system_score_gemma":0.001347214,"threshold_uncertainty_score":0.013529122},"labels":[],"label_agreement":null},{"id":"W4283792637","doi":"10.1609/aaai.v36i11.21591","title":"Multi-Dimension Attention for Multi-Turn Dialog Generation (Student Abstract)","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Dialog box; Dimension (graph theory); Computer science; Generative grammar; Scope (computer science); Generative model; Artificial intelligence; Process (computing); Dialog system; Natural language processing; Semantic interpretation; Speech recognition; Programming language; Mathematics; World Wide Web","score_opus":0.2006134807242427,"score_gpt":0.3464957192817173,"score_spread":0.1458822385574746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283792637","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13969073,0.0008964765,0.8526291,0.0010314597,0.00017395156,0.00008875481,0.00023473821,0.0016606704,0.0035940988],"genre_scores_gemma":[0.9401259,0.00018800731,0.054934356,0.00022458885,0.00007664555,0.000105582796,0.00019039788,0.000091868234,0.004062757],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996296,0.00016660006,0.000012719689,0.00010589067,0.000034102006,0.000051003415],"domain_scores_gemma":[0.99876714,0.0009202151,0.00005998671,0.00007709,0.000115972965,0.00005960333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012495398,0.00054215395,0.0005051647,0.00045990568,0.00037008978,0.0006245738,0.001052411,0.0009907396,0.0031011929],"category_scores_gemma":[0.0027189597,0.00034667063,0.0009106671,0.00038595524,0.0005719589,0.0009600231,0.0012139948,0.0015175627,0.00051545945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035398762,0.00017365116,0.002330375,0.00011124658,0.00013314701,0.00024716728,0.000493772,0.77284515,0.011674315,0.018421084,0.0031575104,0.19005854],"study_design_scores_gemma":[0.000004511222,0.00001658737,0.00019649311,0.0000027209383,0.000007072393,0.000016300053,0.0000050749004,0.9955396,0.0005841331,0.0034603311,0.0001622622,0.0000048816055],"about_ca_topic_score_codex":0.007739891,"about_ca_topic_score_gemma":0.008144091,"teacher_disagreement_score":0.007739891,"about_ca_system_score_codex":0.0009334032,"about_ca_system_score_gemma":0.0005642933,"threshold_uncertainty_score":0.015389681},"labels":[],"label_agreement":null},{"id":"W4283792669","doi":"10.1609/aaai.v36i10.21386","title":"Supervising Model Attention with Human Explanations for Robust Natural Language Inference","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation","keywords":"Computer science; Premise; Inference; Artificial intelligence; Natural (archaeology); Task (project management); Punctuation; Natural language; Natural language processing; Focus (optics); Language model; Machine learning; Linguistics","score_opus":0.11284071439454472,"score_gpt":0.31577590586730336,"score_spread":0.20293519147275862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283792669","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052186023,0.0008409338,0.934971,0.0011284221,0.0001020803,0.00015897777,0.00020472534,0.008078294,0.0023296694],"genre_scores_gemma":[0.7772051,0.0002346528,0.21826543,0.0007872792,0.00014431607,0.00017649797,0.0006178156,0.00046420164,0.0021046796],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974688,0.0012479803,0.00008910045,0.0006963563,0.0003472985,0.0001504321],"domain_scores_gemma":[0.98779035,0.008216796,0.00060627493,0.0022949176,0.0007854195,0.0003062196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046649324,0.0013711948,0.0010502111,0.0009251531,0.00056920137,0.0014656815,0.0027464346,0.0021489463,0.0038510002],"category_scores_gemma":[0.024486259,0.00081616745,0.001004936,0.00070614106,0.0013689655,0.0036124568,0.0026633965,0.004356109,0.0013522324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008731686,0.0006842437,0.010058071,0.00045633086,0.00033990567,0.0003348166,0.0017353774,0.29755968,0.02552943,0.017947556,0.015260734,0.6292207],"study_design_scores_gemma":[0.00005852709,0.000074921285,0.000906023,0.000025306572,0.000043416356,0.00007947381,0.00007258608,0.96445847,0.004664514,0.026933217,0.0026588095,0.000024699373],"about_ca_topic_score_codex":0.0073189256,"about_ca_topic_score_gemma":0.014839629,"teacher_disagreement_score":0.0073189256,"about_ca_system_score_codex":0.0013893249,"about_ca_system_score_gemma":0.0019382917,"threshold_uncertainty_score":0.02467078},"labels":[],"label_agreement":null},{"id":"W4283797513","doi":"10.1609/aaai.v36i10.21428","title":"Unsupervised Sentence Representation via Contrastive Learning with Mixing Negatives","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China; Leverhulme Trust","keywords":"Computer science; Sentence; Artificial intelligence; Representation (politics); Natural language processing; Task (project management); Feature learning; Similarity (geometry); Transfer of learning; Unsupervised learning; Process (computing); Machine learning; Pattern recognition (psychology); Image (mathematics)","score_opus":0.06820739762611006,"score_gpt":0.28246644336630633,"score_spread":0.21425904574019627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283797513","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055493265,0.00029125263,0.93938243,0.0005845635,0.00007747587,0.00012568645,0.00016921802,0.0013985181,0.0024776182],"genre_scores_gemma":[0.7077226,0.00020121127,0.28251043,0.0009395079,0.000201546,0.00041288292,0.0011574548,0.0003273368,0.006526984],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987382,0.00052063836,0.00005683782,0.00039832576,0.00021681703,0.000069215544],"domain_scores_gemma":[0.9969681,0.0017779766,0.00024091275,0.00040915346,0.00047120923,0.0001326247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022315884,0.0012236737,0.0009705824,0.00092686294,0.0006068878,0.0012295364,0.0021208115,0.0013784518,0.002270404],"category_scores_gemma":[0.007983832,0.000490346,0.0008503343,0.0005815771,0.0014437945,0.0032240332,0.002410583,0.0025804404,0.0010604615],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008038707,0.00062505406,0.0050572003,0.00037991727,0.00017462789,0.00045123472,0.0009279402,0.23501262,0.05185749,0.05842924,0.012200104,0.63408077],"study_design_scores_gemma":[0.00002459983,0.0001285058,0.00027733305,0.000013099593,0.000017869044,0.000068102556,0.000033476797,0.9704515,0.0053003114,0.022380568,0.0012897974,0.00001475702],"about_ca_topic_score_codex":0.00079691556,"about_ca_topic_score_gemma":0.0016864054,"teacher_disagreement_score":0.002270404,"about_ca_system_score_codex":0.000754029,"about_ca_system_score_gemma":0.00065350655,"threshold_uncertainty_score":0.0118018985},"labels":[],"label_agreement":null},{"id":"W4283799331","doi":"10.1609/aaai.v36i10.21382","title":"Generation-Focused Table-Based Intermediate Pre-training for Free-Form Question Answering","year":2022,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Question answering; Table (database); Leverage (statistics); Schema (genetic algorithms); Artificial intelligence; Natural language processing; Ask price; Language model; Task (project management); Information retrieval; Data mining","score_opus":0.13912443110804626,"score_gpt":0.3130940979136808,"score_spread":0.17396966680563455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283799331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042030696,0.0015214602,0.9213971,0.0006253482,0.00028275684,0.0005503043,0.0024735574,0.025314191,0.005804642],"genre_scores_gemma":[0.4617139,0.000685788,0.501414,0.0015310217,0.00016355624,0.0012336946,0.022638505,0.0011172212,0.00950226],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990609,0.000310507,0.00005233453,0.00036399742,0.00010943191,0.00010281628],"domain_scores_gemma":[0.9974597,0.0016622392,0.00007644696,0.00039377838,0.00029847416,0.00010934387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013276006,0.0018043006,0.00083460147,0.00086905184,0.0004450843,0.0010952735,0.0024768277,0.0020651969,0.012436083],"category_scores_gemma":[0.005583252,0.0005337048,0.0015498667,0.0007007372,0.000718765,0.0028599342,0.0019690504,0.003175181,0.0048583923],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007307037,0.0006617621,0.004448628,0.0009129642,0.0001571286,0.0005961311,0.0009928201,0.14058612,0.03023801,0.0122644575,0.049612187,0.7587991],"study_design_scores_gemma":[0.00011045881,0.0004400065,0.0010465636,0.00008741228,0.00007867702,0.00028500348,0.00023657246,0.9389282,0.019821847,0.022762913,0.016160822,0.000041474228],"about_ca_topic_score_codex":0.0038163862,"about_ca_topic_score_gemma":0.007199038,"teacher_disagreement_score":0.012436083,"about_ca_system_score_codex":0.00092645147,"about_ca_system_score_gemma":0.0014583437,"threshold_uncertainty_score":0.04160285},"labels":[],"label_agreement":null},{"id":"W4283802444","doi":"10.1609/aaai.v36i10.21325","title":"Predicting Above-Sentence Discourse Structure Using Distant Supervision from Topic Segmentation","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Natural language processing; Computer science; Paragraph; Sentence; Parsing; Automatic summarization; Artificial intelligence; Segmentation; Task (project management); Semantic role labeling; Linguistics","score_opus":0.02308686032587811,"score_gpt":0.27335984647268585,"score_spread":0.25027298614680776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283802444","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42060632,0.004435211,0.5505258,0.0016044637,0.0003894744,0.0002876672,0.0057378625,0.0077811982,0.008632108],"genre_scores_gemma":[0.8553909,0.00070227945,0.12695497,0.00016369266,0.00037153185,0.00020105451,0.010843451,0.0003626316,0.0050094915],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993698,0.00023197004,0.00003270536,0.00025246115,0.000057876816,0.000055089677],"domain_scores_gemma":[0.9964598,0.0023349582,0.00033536062,0.00021601276,0.0005105228,0.00014335244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001176493,0.0010537399,0.0005578607,0.0018246303,0.00073938095,0.0010093042,0.00068447925,0.0010833002,0.0023554415],"category_scores_gemma":[0.00439432,0.0003714876,0.000683717,0.0010884808,0.0003821793,0.0019904512,0.00084513385,0.0014711637,0.0020738088],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018162237,0.0007732751,0.03293547,0.0011623704,0.00029369196,0.00092205155,0.0043361355,0.07319925,0.09828236,0.009527642,0.035387497,0.741364],"study_design_scores_gemma":[0.00007353424,0.0002083014,0.012438216,0.00009880083,0.00014813228,0.00016443462,0.0005780301,0.93486935,0.025656313,0.013934267,0.011784008,0.000046677153],"about_ca_topic_score_codex":0.0039072,"about_ca_topic_score_gemma":0.009540496,"teacher_disagreement_score":0.0039072,"about_ca_system_score_codex":0.00059006445,"about_ca_system_score_gemma":0.0010043522,"threshold_uncertainty_score":0.007879674},"labels":[],"label_agreement":null},{"id":"W4284673342","doi":"10.1145/3477495.3531793","title":"Detecting Frozen Phrases in Open-Domain Question Answering","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Huawei Technologies","keywords":"Computer science; Natural language processing; Question answering; Open domain; Artificial intelligence; Context (archaeology); Domain (mathematical analysis); Task (project management); Information retrieval; Noun phrase; Natural language","score_opus":0.09370089461891895,"score_gpt":0.3472013579749158,"score_spread":0.25350046335599685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284673342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4245999,0.0028764328,0.56512517,0.0011749514,0.00007222747,0.00021064734,0.00090212346,0.002773769,0.002264815],"genre_scores_gemma":[0.86927634,0.0004534091,0.12579574,0.0003934181,0.00012922836,0.000115772185,0.0025187163,0.00016770924,0.0011496287],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968995,0.0017600942,0.00014183018,0.0006124671,0.00042070812,0.00016545046],"domain_scores_gemma":[0.97826433,0.017524632,0.0010910687,0.0016260551,0.0011716605,0.00032223738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050920215,0.00084508036,0.0009943154,0.0019253478,0.00085998565,0.0011695854,0.001668328,0.002562668,0.0011258905],"category_scores_gemma":[0.02709608,0.0006093163,0.00073307037,0.0013933013,0.0016212496,0.0049618096,0.0021315464,0.002138549,0.0010315743],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020395564,0.0007522073,0.041457012,0.001499382,0.0003087157,0.0016692124,0.0075721308,0.14511377,0.08908335,0.023898665,0.013315152,0.6732909],"study_design_scores_gemma":[0.000097620614,0.000527877,0.010505745,0.0000952696,0.00012480069,0.0009870449,0.001225126,0.88032395,0.024549617,0.07549525,0.0059781545,0.000089542766],"about_ca_topic_score_codex":0.0040751314,"about_ca_topic_score_gemma":0.0037275671,"teacher_disagreement_score":0.0050920215,"about_ca_system_score_codex":0.00069335534,"about_ca_system_score_gemma":0.00069028785,"threshold_uncertainty_score":0.026929498},"labels":[],"label_agreement":null},{"id":"W4284682639","doi":"10.1145/3477495.3531749","title":"Document Expansion Baselines and Learned Sparse Lexical Representations for MS MARCO V1 and V2","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Computer science; Weighting; Ranking (information retrieval); Information retrieval; Natural language processing; Question answering; Artificial intelligence; Sequence (biology); Term (time); Language model; Artificial neural network","score_opus":0.12052163994477666,"score_gpt":0.3619357305030352,"score_spread":0.24141409055825852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284682639","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47100493,0.015127661,0.26734692,0.003781539,0.0031843185,0.0040319273,0.07434092,0.09053209,0.07064975],"genre_scores_gemma":[0.46892262,0.0016483358,0.3058983,0.0012959033,0.0005076048,0.003468305,0.19358882,0.0044011874,0.020268938],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949443,0.0017122373,0.00039366714,0.0011734955,0.0013651315,0.00041115115],"domain_scores_gemma":[0.9931986,0.002066365,0.0003017843,0.0023793154,0.0017022244,0.00035170678],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077228695,0.0024499332,0.0015346812,0.00425943,0.0016888373,0.0025842453,0.004069519,0.0019874838,0.00724272],"category_scores_gemma":[0.02344943,0.0007176053,0.001585168,0.0034220517,0.0010539056,0.0046404493,0.0031603482,0.0032780997,0.0058149816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031597146,0.0024993801,0.009409162,0.0019104274,0.0007441691,0.00044236204,0.0004646883,0.108134486,0.025809979,0.007744473,0.25543094,0.5842503],"study_design_scores_gemma":[0.0010478152,0.0020864264,0.011999246,0.00023414212,0.0003681083,0.0005968324,0.0006397059,0.8538268,0.048597038,0.010521751,0.06980337,0.00027877977],"about_ca_topic_score_codex":0.02789957,"about_ca_topic_score_gemma":0.043617215,"teacher_disagreement_score":0.02789957,"about_ca_system_score_codex":0.0027320916,"about_ca_system_score_gemma":0.0022168346,"threshold_uncertainty_score":0.05547434},"labels":[],"label_agreement":null},{"id":"W4284686595","doi":"10.1145/3477495.3531717","title":"Another Look at Information Retrieval as Statistical Translation","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Ranking (information retrieval); Simple (philosophy); Information retrieval; Transformer; Statistical model; Replication (statistics); Machine learning; Natural language processing","score_opus":0.08262252473725214,"score_gpt":0.32561120261350107,"score_spread":0.24298867787624895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284686595","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012131271,0.10536185,0.6365465,0.1516387,0.00453037,0.00014995213,0.0011287994,0.0019794456,0.08653304],"genre_scores_gemma":[0.476286,0.06418786,0.29960322,0.055894263,0.021177875,0.0004749868,0.001928055,0.0019401965,0.07850763],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99686867,0.0014795637,0.0001429057,0.00065525126,0.0006899333,0.00016355324],"domain_scores_gemma":[0.9917663,0.005171953,0.00034705814,0.0016031316,0.0008897721,0.0002217392],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004636724,0.0012786065,0.0014750375,0.004568147,0.001386359,0.0075224442,0.001799544,0.0040731947,0.009380486],"category_scores_gemma":[0.017265199,0.0007565608,0.0017859354,0.005036436,0.008723887,0.027369227,0.0030806342,0.0075956257,0.003944972],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000104354396,0.00008061753,0.00048342385,0.00033329488,0.00007291374,0.00007091982,0.00039752072,0.0050085457,0.0012261152,0.8952377,0.027822854,0.069161765],"study_design_scores_gemma":[0.000060925446,0.00018418518,0.00061232096,0.00025167782,0.000052087602,0.0003098978,0.00023242878,0.03981252,0.002048861,0.8284407,0.1279055,0.000088882596],"about_ca_topic_score_codex":0.00408143,"about_ca_topic_score_gemma":0.002224917,"teacher_disagreement_score":0.009380486,"about_ca_system_score_codex":0.0034230244,"about_ca_system_score_gemma":0.0014381224,"threshold_uncertainty_score":0.03138089},"labels":[],"label_agreement":null},{"id":"W4284689311","doi":"10.1145/3477495.3531986","title":"H-ERNIE","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wilfrid Laurier University","funders":"","keywords":"Computer science; Relevance (law); Ambiguity; Focus (optics); Matching (statistics); Mores; Language model; Natural language; Information retrieval; Natural language processing; Artificial intelligence; Programming language","score_opus":0.0967995303531606,"score_gpt":0.3273152172855074,"score_spread":0.23051568693234678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284689311","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023338161,0.0066096853,0.133354,0.013908362,0.011165333,0.00064394186,0.011173706,0.020041922,0.7797649],"genre_scores_gemma":[0.1352332,0.002999102,0.06315598,0.0041103433,0.0017225959,0.00033414972,0.017872583,0.0022687772,0.7723034],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992072,0.00012820728,0.000028490327,0.00025289136,0.0002550309,0.00012827136],"domain_scores_gemma":[0.9989492,0.00016636428,0.000043683573,0.00025320443,0.00041286598,0.0001747038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012906989,0.00095753063,0.00080911524,0.0010331667,0.0012150569,0.0019455225,0.00094776304,0.0012206946,0.17946337],"category_scores_gemma":[0.003034172,0.00031579676,0.00045431676,0.00082566694,0.0004878727,0.002443996,0.002222696,0.0016685677,0.11807754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005931939,0.00020153633,0.0018724573,0.00036006802,0.000055672743,0.0003317222,0.00016956241,0.002713128,0.008527428,0.03959766,0.39607817,0.54949945],"study_design_scores_gemma":[0.00009483234,0.00020566076,0.0014838872,0.000111551766,0.000029549541,0.00042192813,0.00012232186,0.016918097,0.009201166,0.016693885,0.95464826,0.00006885948],"about_ca_topic_score_codex":0.0031485674,"about_ca_topic_score_gemma":0.0044412203,"teacher_disagreement_score":0.17946337,"about_ca_system_score_codex":0.0010884436,"about_ca_system_score_gemma":0.0019857055,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4284706463","doi":"10.1145/3477495.3531898","title":"Inconsistent Ranking Assumptions in Medical Search and Their Downstream Consequences","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"National Science Foundation","keywords":"Downstream (manufacturing); Ranking (information retrieval); Relevance (law); Computer science; Point (geometry); Machine learning; Confidence interval; Artificial intelligence; Information retrieval; Data mining; Statistics; Mathematics; Engineering","score_opus":0.09496922730706166,"score_gpt":0.3320899544838965,"score_spread":0.23712072717683486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284706463","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19155662,0.0046299486,0.6997485,0.059162434,0.00043701418,0.0003450756,0.0016879593,0.0014248026,0.04100768],"genre_scores_gemma":[0.90945035,0.00142146,0.07711493,0.003829171,0.0009613264,0.00037797014,0.0011553048,0.0003761429,0.0053132665],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9723092,0.013939437,0.0018391978,0.003971446,0.006592947,0.0013477146],"domain_scores_gemma":[0.74924105,0.19980404,0.015531764,0.021261808,0.012186899,0.0019743869],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04506032,0.00093312707,0.0034342422,0.00348263,0.0030270405,0.005646696,0.0056747356,0.0069358493,0.007723524],"category_scores_gemma":[0.28490102,0.0021721728,0.0016163451,0.0039110174,0.008157579,0.018052174,0.0062841433,0.010909204,0.0023638997],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071412924,0.00030821608,0.017564563,0.00047219684,0.00017560166,0.0013279484,0.0012812195,0.06489401,0.0011189136,0.82474935,0.013405321,0.07398855],"study_design_scores_gemma":[0.000096310985,0.000059765734,0.0018950611,0.00007468964,0.000028585342,0.00036868264,0.0001408443,0.11094933,0.00065428915,0.8844479,0.0012205157,0.00006398873],"about_ca_topic_score_codex":0.004908707,"about_ca_topic_score_gemma":0.0031283605,"teacher_disagreement_score":0.04506032,"about_ca_system_score_codex":0.0030006561,"about_ca_system_score_gemma":0.0020980062,"threshold_uncertainty_score":0.23830462},"labels":[],"label_agreement":null},{"id":"W4284708021","doi":"10.1145/3477495.3531853","title":"Neural Query Synthesis and Domain-Specific Ranking Templates for Multi-Stage Clinical Trial Matching","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Ranking (information retrieval); Computer science; Matching (statistics); Pipeline (software); Pointwise; Relevance (law); Sentence; Domain (mathematical analysis); Artificial intelligence; Information retrieval; Data mining","score_opus":0.30087903734505234,"score_gpt":0.4115527108054265,"score_spread":0.11067367346037416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284708021","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015919147,0.00099313,0.9577998,0.00062399317,0.00019489923,0.00069867657,0.00092202186,0.020643497,0.0022049053],"genre_scores_gemma":[0.22274691,0.0003757481,0.76595443,0.00077361334,0.00027279794,0.0006138759,0.0032183332,0.0006305698,0.0054136375],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99454737,0.0023162307,0.0005960128,0.0010056513,0.0013072011,0.0002275178],"domain_scores_gemma":[0.991521,0.0043333424,0.00061745296,0.0016385881,0.0016090879,0.00028051337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007199746,0.0011440231,0.0012672306,0.0023736751,0.0005093128,0.001539594,0.0027065754,0.0016869386,0.00673253],"category_scores_gemma":[0.01810601,0.0005182148,0.0011043177,0.0014631212,0.0005860567,0.0026803052,0.0016135505,0.0016011533,0.0046210224],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011289776,0.00059634435,0.002155272,0.0007714054,0.0001904403,0.0003674353,0.0003222582,0.057963125,0.037132785,0.008002235,0.028483443,0.86288625],"study_design_scores_gemma":[0.00022552728,0.00067830586,0.0010533516,0.000049979353,0.00011439896,0.0004807267,0.00010455605,0.92675817,0.0394945,0.015324639,0.015632482,0.00008340857],"about_ca_topic_score_codex":0.0037768392,"about_ca_topic_score_gemma":0.006130682,"teacher_disagreement_score":0.007199746,"about_ca_system_score_codex":0.0013596851,"about_ca_system_score_gemma":0.0023209145,"threshold_uncertainty_score":0.03807634},"labels":[],"label_agreement":null},{"id":"W4284709654","doi":"10.1145/3510003.3510159","title":"CLEAR","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 44th International Conference on Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Task (project management); Natural language processing; Artificial intelligence; Word (group theory); Information retrieval; Embedding; Semantics (computer science); Programming language; Mathematics","score_opus":0.02487960478481576,"score_gpt":0.22987186236868118,"score_spread":0.2049922575838654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4284709654","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014455273,0.0062199133,0.4401936,0.004653733,0.0026834353,0.0010531719,0.039009,0.12514542,0.36658654],"genre_scores_gemma":[0.12521695,0.0051580835,0.26723787,0.003846587,0.0010379626,0.0007515481,0.12765297,0.02170447,0.44739357],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975248,0.00031258428,0.00019325259,0.000733267,0.0010005955,0.00023560897],"domain_scores_gemma":[0.99562705,0.0005993117,0.00028156143,0.0017055698,0.0015943267,0.00019215237],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016898963,0.0016324151,0.0009828237,0.0030805082,0.0014880387,0.0043710913,0.0020604117,0.0016720807,0.12776339],"category_scores_gemma":[0.007465617,0.0007992721,0.0012783081,0.0026979272,0.0005984146,0.008078856,0.002925541,0.0017125119,0.16235392],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003591476,0.0001250791,0.0032143595,0.0009726884,0.00008075906,0.00033294107,0.00022432432,0.0017868285,0.0076345,0.037458234,0.4081528,0.5396584],"study_design_scores_gemma":[0.00003799892,0.00007034931,0.0028885708,0.00018950747,0.000054080905,0.00067241356,0.0001918831,0.014376384,0.008497684,0.02000208,0.9529409,0.000078076395],"about_ca_topic_score_codex":0.0050761895,"about_ca_topic_score_gemma":0.0067111193,"teacher_disagreement_score":0.12776339,"about_ca_system_score_codex":0.00075602817,"about_ca_system_score_gemma":0.0016005567,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4285009115","doi":"10.22215/etd/2022-15107","title":"Relation Mapping For Question Answering Over Knowledge Graphs Using Large Corpus Of Free Text","year":2022,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Question answering; Knowledge graph; Relation (database); Natural language; Natural language processing; SPARQL; Artificial intelligence; Information retrieval; Graph; Context (archaeology); RDF; Semantic Web; Theoretical computer science; Data mining","score_opus":0.030199614223749248,"score_gpt":0.30425153730216664,"score_spread":0.2740519230784174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285009115","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011153299,0.0007176184,0.9697963,0.0004533226,0.000046777637,0.00033611926,0.0028586637,0.012590566,0.0020472566],"genre_scores_gemma":[0.1391654,0.00066487124,0.8433292,0.000233481,0.000068113004,0.0006102848,0.012954528,0.00081488764,0.0021592814],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747795,0.00090097374,0.00016156111,0.0008468774,0.00051753817,0.00009513371],"domain_scores_gemma":[0.9941292,0.0043855836,0.00025779544,0.00074895984,0.00036886946,0.00010963404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019044707,0.0012670906,0.0009763304,0.006361397,0.0011463774,0.0019178953,0.0018948322,0.0013757227,0.00568505],"category_scores_gemma":[0.011806304,0.0006804888,0.0023603602,0.0043943594,0.00091255107,0.008003835,0.0028458189,0.0014163238,0.0025021369],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036469573,0.0005275392,0.003383492,0.002163259,0.0002969058,0.0008065082,0.00286778,0.082146004,0.01874563,0.08390839,0.030588772,0.7742011],"study_design_scores_gemma":[0.00007456492,0.000119136224,0.0016299031,0.0001508015,0.00011036778,0.00042636177,0.0009744884,0.675697,0.011364705,0.26747262,0.041916355,0.000063649575],"about_ca_topic_score_codex":0.0061635477,"about_ca_topic_score_gemma":0.0098324735,"teacher_disagreement_score":0.006361397,"about_ca_system_score_codex":0.001236627,"about_ca_system_score_gemma":0.0012587095,"threshold_uncertainty_score":0.019018412},"labels":[],"label_agreement":null},{"id":"W4285045211","doi":"10.1007/s10796-022-10295-0","title":"Knowledge Graph and Deep Learning-based Text-to-GraphQL Model for Intelligent Medical Consultation Chatbot","year":2022,"lang":"en","type":"article","venue":"Information Systems Frontiers","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Adapter (computing); Chatbot; Natural language processing; Artificial intelligence; Parsing; Programming language; Information retrieval; Database","score_opus":0.017524946325652855,"score_gpt":0.2484159739786928,"score_spread":0.23089102765303995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285045211","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0870932,0.0010963109,0.876839,0.0020212,0.00022584971,0.00043390176,0.0039695,0.01691999,0.0114010535],"genre_scores_gemma":[0.8300214,0.0005687012,0.13333482,0.0011214743,0.00008760963,0.0006994384,0.0071827048,0.00034997493,0.02663384],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99961,0.000050190298,0.000020590693,0.00019380215,0.000058864036,0.00006655306],"domain_scores_gemma":[0.99971145,0.00010995062,0.000024735293,0.00003523851,0.00008087813,0.000037651735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00043480727,0.0010735164,0.0006053607,0.0008628651,0.0005631751,0.000989009,0.0023207588,0.0016200066,0.009190645],"category_scores_gemma":[0.0013092185,0.00038948873,0.0010964439,0.0006177229,0.0004371308,0.0019081287,0.0010796444,0.0014345722,0.0016760493],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000605273,0.0007932154,0.0044390433,0.00040274256,0.00016104255,0.000699476,0.00045522346,0.6001549,0.009133002,0.019931277,0.029754398,0.3334704],"study_design_scores_gemma":[0.000010440537,0.00002628067,0.00020738464,0.0000051197608,0.000014904847,0.000015956337,0.000021431257,0.9946479,0.00065115053,0.0034259006,0.000967148,0.000006364321],"about_ca_topic_score_codex":0.03846155,"about_ca_topic_score_gemma":0.034674954,"teacher_disagreement_score":0.03846155,"about_ca_system_score_codex":0.0021834446,"about_ca_system_score_gemma":0.0019890252,"threshold_uncertainty_score":0.07647538},"labels":[],"label_agreement":null},{"id":"W4285119752","doi":"10.18653/v1/2022.bionlp-1.33","title":"Doctor XAvIer: Explainable Diagnosis on Physician-Patient Dialogues and XAI Evaluation","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Attribution; Feature (linguistics); Metric (unit); Plot (graphics); Artificial intelligence; Computer science; Pattern recognition (psychology); F1 score; Natural language processing; Mathematics; Statistics; Psychology; Linguistics; Social psychology; Philosophy; Engineering","score_opus":0.046959093352888046,"score_gpt":0.2631326237226032,"score_spread":0.21617353036971515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285119752","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.528801,0.0046722586,0.30042678,0.0037821103,0.0009966106,0.0020748805,0.01428038,0.12547866,0.019487191],"genre_scores_gemma":[0.8039474,0.00079098786,0.1727886,0.00054223964,0.00019026807,0.0004470377,0.012936496,0.0009807991,0.007376162],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960114,0.0021486622,0.00021373185,0.000702498,0.0007393705,0.00018436385],"domain_scores_gemma":[0.9828348,0.013704856,0.00059889985,0.00090879836,0.0012960291,0.0006566074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067066834,0.0013991552,0.00071370014,0.0017851288,0.00049922994,0.0017322733,0.0014238606,0.0017973979,0.0067469003],"category_scores_gemma":[0.024502587,0.00035525332,0.0006649665,0.00066472887,0.0004653235,0.0014329616,0.0015162238,0.0014472565,0.002171493],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006970185,0.0017586225,0.05155579,0.0022744932,0.0006436084,0.0021991904,0.0054777768,0.09829876,0.036799025,0.005364184,0.06411728,0.72454107],"study_design_scores_gemma":[0.0005936423,0.0017692893,0.02588737,0.00020575765,0.00018186314,0.001346999,0.0010815935,0.89149475,0.039066575,0.0037952708,0.03437035,0.00020653597],"about_ca_topic_score_codex":0.00553627,"about_ca_topic_score_gemma":0.0053783045,"teacher_disagreement_score":0.0067469003,"about_ca_system_score_codex":0.001350142,"about_ca_system_score_gemma":0.0011869898,"threshold_uncertainty_score":0.035468757},"labels":[],"label_agreement":null},{"id":"W4285143696","doi":"10.18653/v1/2022.in2writing-1","title":"Proceedings of the First Workshop on Intelligent and Interactive Writing Assistants (In2Writing 2022)","year":2022,"lang":"en","type":"paratext","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"National Science Foundation","keywords":"Computer science; Human–computer interaction; Multimedia","score_opus":0.02983464763860292,"score_gpt":0.27783615263109507,"score_spread":0.24800150499249216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285143696","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033470083,0.048786536,0.41799763,0.027474375,0.06729581,0.0017732331,0.0060909037,0.008531859,0.38857964],"genre_scores_gemma":[0.070174925,0.018693434,0.19659685,0.0039021454,0.00844744,0.0012760333,0.014543984,0.0040142178,0.68235105],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99619603,0.0017253989,0.00021332546,0.00073904055,0.00077794935,0.0003483496],"domain_scores_gemma":[0.99428046,0.0019782705,0.00015604617,0.00084739714,0.0015462177,0.0011915758],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006350462,0.0014615435,0.001302413,0.0017479407,0.0016901409,0.0068783686,0.0026479345,0.0026611316,0.07417146],"category_scores_gemma":[0.008714819,0.000637573,0.0010016151,0.0018305881,0.0012075694,0.0066513973,0.004565949,0.003117703,0.029864717],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074703566,0.00046469553,0.0007472187,0.00090664264,0.00007054545,0.00038528178,0.001958635,0.0016830872,0.006281047,0.023648538,0.5194068,0.44370055],"study_design_scores_gemma":[0.000039260736,0.00010904845,0.00060892163,0.00023396051,0.000028939086,0.0002412146,0.00052122434,0.0031983368,0.0026105607,0.007199248,0.9851812,0.000028113185],"about_ca_topic_score_codex":0.0027779671,"about_ca_topic_score_gemma":0.0053490982,"teacher_disagreement_score":0.07417146,"about_ca_system_score_codex":0.0016275576,"about_ca_system_score_gemma":0.002561257,"threshold_uncertainty_score":0.24812824},"labels":[],"label_agreement":null},{"id":"W4285148079","doi":"10.18653/v1/2022.findings-acl.293","title":"Local Structure Matters Most: Perturbation Study in NLU","year":2022,"lang":"en","type":"article","venue":"Findings of the Association for Computational Linguistics: ACL 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Canada First Research Excellence Fund; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Perturbation (astronomy); Word order; Phenomenon; Artificial neural network; Invariant (physics); Artificial intelligence; Natural language processing; Mathematics; Physics","score_opus":0.010150799814172976,"score_gpt":0.24417974456153907,"score_spread":0.2340289447473661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285148079","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9346654,0.00047688815,0.06195209,0.0005671018,0.000027730279,0.000046848007,0.00007889848,0.0003223946,0.0018626552],"genre_scores_gemma":[0.9949645,0.00006364839,0.004639853,0.000049752696,0.000012910659,0.000019638363,0.000051909425,0.00004494576,0.00015274045],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99860185,0.00091968477,0.000047899954,0.00023708747,0.00014227604,0.000051316518],"domain_scores_gemma":[0.9796521,0.016794078,0.0008062698,0.0019396228,0.00049690896,0.00031100394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002466983,0.00034083024,0.00057369564,0.00044032556,0.00041083485,0.0009341705,0.00060907303,0.0007648727,0.0010041981],"category_scores_gemma":[0.033530142,0.0002995041,0.00030692236,0.00042781266,0.0012964723,0.0027097592,0.00093932706,0.0016189405,0.00018933485],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027328664,0.0010644355,0.0444022,0.00053365837,0.0004686507,0.0006758727,0.0040427293,0.6395189,0.14533497,0.03391996,0.0022522218,0.12505351],"study_design_scores_gemma":[0.00004116764,0.00029010462,0.014469592,0.000021912398,0.000048947168,0.000108375025,0.000371936,0.9240426,0.020874677,0.03900556,0.0006899219,0.000035261513],"about_ca_topic_score_codex":0.0015410637,"about_ca_topic_score_gemma":0.0011587379,"teacher_disagreement_score":0.002466983,"about_ca_system_score_codex":0.00064279296,"about_ca_system_score_gemma":0.00028849972,"threshold_uncertainty_score":0.013046801},"labels":[],"label_agreement":null},{"id":"W4285155190","doi":"10.18653/v1/2022.findings-acl.168","title":"Question Generation for Reading Comprehension Assessment by Modeling How and What to Ask","year":2022,"lang":"en","type":"article","venue":"Findings of the Association for Computational Linguistics: ACL 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Reading comprehension; Computer science; Comprehension; Reading (process); Focus (optics); Artificial intelligence; Natural language processing; Mathematics education; Linguistics; Psychology","score_opus":0.02239636680911099,"score_gpt":0.2829373594837503,"score_spread":0.2605409926746393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285155190","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22219907,0.0064944634,0.60186166,0.004943758,0.0007422758,0.0030653747,0.08773612,0.05824472,0.014712581],"genre_scores_gemma":[0.39236644,0.0007003239,0.41570556,0.0011673897,0.000195863,0.002426306,0.18120879,0.00087994774,0.0053494032],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968746,0.0016660356,0.00022254042,0.000853925,0.0002914026,0.00009143002],"domain_scores_gemma":[0.9868009,0.009019402,0.0005819158,0.0019199513,0.0013731784,0.0003045892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032771386,0.0019598466,0.00061732094,0.002861357,0.0006599366,0.0016753619,0.0027680222,0.0030816114,0.005406877],"category_scores_gemma":[0.022015857,0.00044099687,0.0018138688,0.0015028715,0.0006542341,0.004015435,0.0018446759,0.003114311,0.0051353485],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013665336,0.0023153657,0.06082597,0.0033944335,0.0006029955,0.0006955085,0.0031235549,0.07085049,0.022492453,0.009141075,0.13241144,0.69278026],"study_design_scores_gemma":[0.0003295262,0.00075559574,0.023405893,0.00032725005,0.00026007328,0.0006765441,0.0009786909,0.8290964,0.023978934,0.031021897,0.08902275,0.00014643806],"about_ca_topic_score_codex":0.0063370056,"about_ca_topic_score_gemma":0.014286781,"teacher_disagreement_score":0.0063370056,"about_ca_system_score_codex":0.0014490603,"about_ca_system_score_gemma":0.001316939,"threshold_uncertainty_score":0.018087864},"labels":[],"label_agreement":null},{"id":"W4285171765","doi":"10.18653/v1/2022.acl-long.360","title":"PRIMERA: Pyramid-based Masked Sentence Pre-training for Multi-document Summarization","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":119,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Encoder; Transformer; Focus (optics); Artificial intelligence; Pyramid (geometry); Sentence; Multi-document summarization; Training set; Natural language processing; Information retrieval; Pattern recognition (psychology)","score_opus":0.019584428927741152,"score_gpt":0.26176870111486816,"score_spread":0.24218427218712701,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285171765","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064500426,0.00072154315,0.9653701,0.00022904991,0.0002426227,0.00022736847,0.0015168199,0.023636324,0.0016061459],"genre_scores_gemma":[0.11426127,0.00059725455,0.8589962,0.00062781124,0.00027811917,0.0008222105,0.012258807,0.0016557733,0.010502557],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931586,0.00017459484,0.000046732668,0.00025010033,0.00013903051,0.00007376316],"domain_scores_gemma":[0.9985575,0.00057585037,0.000092015274,0.0002900863,0.00039444325,0.00009000177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011812026,0.0018559555,0.0010434048,0.0011839415,0.00057959626,0.0010881942,0.0025805617,0.0014687666,0.008744086],"category_scores_gemma":[0.0037145761,0.00076847203,0.0011544313,0.0011790629,0.00049584993,0.002690153,0.0017868477,0.0030370508,0.008520286],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040952733,0.00029310235,0.0008091203,0.0005113073,0.00018941029,0.00020228414,0.00037136883,0.0450198,0.04758563,0.004382497,0.05016866,0.8500573],"study_design_scores_gemma":[0.00007767155,0.00044778344,0.0008359156,0.000058921207,0.00011257172,0.0002035712,0.00014195808,0.91199636,0.046347186,0.012771635,0.026945269,0.00006107013],"about_ca_topic_score_codex":0.0042333878,"about_ca_topic_score_gemma":0.010757701,"teacher_disagreement_score":0.008744086,"about_ca_system_score_codex":0.0007308609,"about_ca_system_score_gemma":0.0018343026,"threshold_uncertainty_score":0.029251873},"labels":[],"label_agreement":null},{"id":"W4285226597","doi":"10.18653/v1/2022.acl-long.92","title":"WatClaimCheck: A new Dataset for Claim Entailment and Inference","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Waterloo","funders":"Vector Institute; Natural Sciences and Engineering Research Council of Canada; Government of Canada; Canadian Institute for Advanced Research","keywords":"Premise; Inference; Computer science; Identification (biology); Task (project management); Logical consequence; Textual entailment; Information retrieval; Quality (philosophy); Data science; Artificial intelligence; Natural language processing; Epistemology; Engineering","score_opus":0.014398754585220922,"score_gpt":0.2566326457056536,"score_spread":0.24223389112043267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285226597","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028156398,0.0028475598,0.016713701,0.0018596449,0.00046907397,0.00077461085,0.92844474,0.00833891,0.012395392],"genre_scores_gemma":[0.015920663,0.00041300076,0.023763075,0.00024529872,0.00012656143,0.00062624976,0.9561655,0.0003560559,0.002383607],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9966419,0.0005895108,0.0007153586,0.0008979875,0.0009614221,0.00019391009],"domain_scores_gemma":[0.9887754,0.0051784036,0.0013522407,0.0021887629,0.001884014,0.0006212268],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002182353,0.0019569313,0.0009762843,0.01099814,0.0022099486,0.0030267325,0.0037848805,0.0045959507,0.01786221],"category_scores_gemma":[0.017284732,0.000726083,0.0018435786,0.007404202,0.0008713146,0.0055216006,0.0037500744,0.0025940211,0.0114752585],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005693147,0.000668478,0.010757425,0.0046319473,0.00025339652,0.0014858269,0.0009980442,0.0039932807,0.005826803,0.013681091,0.88326466,0.07386983],"study_design_scores_gemma":[0.0005486725,0.00018069055,0.02282144,0.0006721976,0.00019042329,0.0014626301,0.0011890593,0.031768087,0.009332157,0.016020905,0.9156517,0.00016216532],"about_ca_topic_score_codex":0.016193403,"about_ca_topic_score_gemma":0.038824786,"teacher_disagreement_score":0.01786221,"about_ca_system_score_codex":0.0022769233,"about_ca_system_score_gemma":0.0036625164,"threshold_uncertainty_score":0.059755027},"labels":[],"label_agreement":null},{"id":"W4285263440","doi":"10.18653/v1/2022.acl-long.129","title":"Probing as Quantifying Inductive Bias","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Max Planck ETH Center for Learning Systems; McGill University; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; National Science Foundation","keywords":"Inductive bias; Computer science; ENCODE; Variety (cybernetics); Task (project management); Bayesian probability; Artificial intelligence; Natural language processing; Range (aeronautics); Position paper; Multi-task learning; Machine learning","score_opus":0.02919653589517536,"score_gpt":0.2628133028168727,"score_spread":0.23361676692169736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285263440","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039041225,0.0007167027,0.9510226,0.0021909245,0.00007060238,0.00011041684,0.00026418752,0.0005803875,0.0060030823],"genre_scores_gemma":[0.7392645,0.0006603995,0.25472918,0.0014092815,0.00032209742,0.00062171463,0.0006805407,0.00043101484,0.0018812224],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98250926,0.010662879,0.0006788668,0.0023528766,0.003086391,0.0007098279],"domain_scores_gemma":[0.8467169,0.12948066,0.006136586,0.0129293455,0.0036130869,0.0011234609],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024869582,0.0013105046,0.001469648,0.002431373,0.0015214636,0.0031906385,0.0027384704,0.0028875484,0.0033141344],"category_scores_gemma":[0.15873833,0.0010683606,0.0010356244,0.002782855,0.006869228,0.012994106,0.008224417,0.0054805013,0.00067556533],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007434618,0.00018778253,0.015699469,0.0006290749,0.00026610517,0.0001876946,0.0023657603,0.12254814,0.008044601,0.6503749,0.0036747805,0.19527826],"study_design_scores_gemma":[0.000043124608,0.00010672109,0.0017849425,0.000080174934,0.00006201019,0.00014656613,0.00017176317,0.16833833,0.0051840865,0.82019126,0.0038353463,0.000055611126],"about_ca_topic_score_codex":0.0012781036,"about_ca_topic_score_gemma":0.0012313278,"teacher_disagreement_score":0.024869582,"about_ca_system_score_codex":0.002293624,"about_ca_system_score_gemma":0.0019496536,"threshold_uncertainty_score":0.13152444},"labels":[],"label_agreement":null},{"id":"W4285264466","doi":"10.2139/ssrn.4101003","title":"Regularizing Deep Text Models by Encouraging Competition","year":2022,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Competition (biology); Computer science; Psychology; Data science","score_opus":0.008238387120771863,"score_gpt":0.20618098105084953,"score_spread":0.19794259393007768,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285264466","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18012996,0.0024843276,0.79452187,0.005140989,0.0010955511,0.00026269438,0.0016329236,0.0063432334,0.008388477],"genre_scores_gemma":[0.89936477,0.0005843499,0.07442966,0.0014043228,0.0011992512,0.00040905245,0.003937785,0.0012732623,0.017397568],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99613816,0.0017078079,0.00020921737,0.0010989406,0.00047899008,0.00036687427],"domain_scores_gemma":[0.97266114,0.021152876,0.0009541205,0.00228565,0.0018718552,0.0010742345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065482706,0.002334172,0.0024275237,0.0018453213,0.0013526002,0.0029659844,0.0036629497,0.0048136595,0.0069957217],"category_scores_gemma":[0.03326687,0.0014251866,0.0016362009,0.0017794534,0.0015211051,0.0076444596,0.005125025,0.0071199285,0.003768848],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024953964,0.0011386914,0.008576086,0.00073729386,0.00048633374,0.00037519506,0.00073361647,0.6501724,0.013242023,0.03927972,0.047620703,0.23514251],"study_design_scores_gemma":[0.00005515714,0.000092852075,0.0002211459,0.000016675567,0.000041455387,0.000030830688,0.000026686548,0.97922486,0.0009975563,0.018209718,0.0010703714,0.000012635019],"about_ca_topic_score_codex":0.0037863425,"about_ca_topic_score_gemma":0.008667234,"teacher_disagreement_score":0.0069957217,"about_ca_system_score_codex":0.0013572844,"about_ca_system_score_gemma":0.0018109418,"threshold_uncertainty_score":0.034630954},"labels":[],"label_agreement":null},{"id":"W4285272223","doi":"10.18653/v1/2022.acl-short.10","title":"Automatic Detection of Entity-Manipulated Text using Factual Knowledge","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Compute Canada; Advanced Micro Devices","keywords":"Computer science; Convolutional neural network; Exploit; Focus (optics); Task (project management); Artificial intelligence; Information retrieval; Code (set theory); Graph; Knowledge graph; Natural language processing; Theoretical computer science; Computer security","score_opus":0.04609297090021416,"score_gpt":0.27107876610736004,"score_spread":0.22498579520714587,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285272223","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41882333,0.010175457,0.51742464,0.0019950473,0.0012056358,0.0005373782,0.014385486,0.012120347,0.023332613],"genre_scores_gemma":[0.7834646,0.0017176177,0.18454577,0.00037992548,0.0006475045,0.0001532703,0.020025536,0.00040593362,0.008659798],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985764,0.00025858617,0.00011442248,0.00049307,0.00045695202,0.00010051065],"domain_scores_gemma":[0.9924273,0.0039903815,0.001401945,0.0007928706,0.0011982218,0.00018927161],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001308903,0.00089295546,0.00048095259,0.00506356,0.00068980287,0.0019377844,0.0014112708,0.0013689675,0.0026509583],"category_scores_gemma":[0.007827323,0.00043262303,0.00053634593,0.0023200768,0.00067950675,0.0053405142,0.0011353764,0.0012380707,0.0014963092],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012251082,0.0003911717,0.05030691,0.0022472958,0.00044370117,0.0056993742,0.0023416867,0.021989217,0.11586896,0.034595646,0.053283557,0.7116074],"study_design_scores_gemma":[0.00006794419,0.00018369897,0.034209464,0.000306929,0.00037356745,0.0039722235,0.0009158948,0.6182604,0.18418595,0.029897109,0.1274844,0.00014250551],"about_ca_topic_score_codex":0.0018492662,"about_ca_topic_score_gemma":0.0032246246,"teacher_disagreement_score":0.00506356,"about_ca_system_score_codex":0.00079308497,"about_ca_system_score_gemma":0.0005642145,"threshold_uncertainty_score":0.008868277},"labels":[],"label_agreement":null},{"id":"W4285276820","doi":"10.18653/v1/2022.bionlp-1.23","title":"BioCite: A Deep Learning-based Citation Linkage Framework for Biomedical Research Articles","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Citation; Sentence; Task (project management); Linkage (software); Information retrieval; Natural language processing; Tree (set theory); Artificial intelligence; Citation analysis; Data science; World Wide Web","score_opus":0.09386321376941074,"score_gpt":0.3667553531999227,"score_spread":0.27289213943051194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285276820","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020144526,0.0044492166,0.9443894,0.0027970804,0.0005063054,0.00032157573,0.00944093,0.010976755,0.0069741667],"genre_scores_gemma":[0.24453881,0.004970783,0.69264275,0.0009541751,0.0009969465,0.0013255762,0.026321268,0.0012182503,0.02703146],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99909973,0.00021661482,0.000083828585,0.00019347429,0.0003290057,0.00007740946],"domain_scores_gemma":[0.9976507,0.0009958933,0.00029777782,0.00018697565,0.000705348,0.00016332945],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.002238803,0.00092298456,0.0008790796,0.00807964,0.0012268206,0.002194171,0.002833753,0.0019511981,0.004744675],"category_scores_gemma":[0.007925395,0.00045826813,0.0015192203,0.0073151845,0.0005966159,0.0042826654,0.0025090172,0.0022491983,0.0020161383],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035027985,0.0005588884,0.005819676,0.0010985635,0.00042211477,0.0004674756,0.0006896281,0.13807587,0.005716855,0.10682726,0.07922711,0.6607463],"study_design_scores_gemma":[0.000044933207,0.00008251078,0.0010005779,0.00010450712,0.00010077772,0.00014107728,0.00010707045,0.85516477,0.0035598031,0.094409615,0.045235865,0.000048540656],"about_ca_topic_score_codex":0.010700581,"about_ca_topic_score_gemma":0.0277731,"teacher_disagreement_score":0.99192035,"about_ca_system_score_codex":0.0025123812,"about_ca_system_score_gemma":0.0036230246,"threshold_uncertainty_score":0.021276593},"labels":[],"label_agreement":null},{"id":"W4285286742","doi":"10.32604/cmc.2022.027236","title":"Automating Transfer Credit Assessment-A Natural Language Processing-Based Approach","year":2022,"lang":"en","type":"article","venue":"Computers, materials & continua/Computers, materials & continua (Print)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Natural language processing; Transfer (computing); Artificial intelligence","score_opus":0.011567822472052763,"score_gpt":0.2432532928695259,"score_spread":0.23168547039747314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285286742","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062333543,0.00031092702,0.91972387,0.0010136281,0.00011399173,0.0010450295,0.0018223813,0.0066900253,0.0069465814],"genre_scores_gemma":[0.4846809,0.000289304,0.50570637,0.00024178643,0.00007130758,0.0005745569,0.005136306,0.00024375507,0.0030556552],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995926,0.0012053939,0.00037678602,0.0011535833,0.0010965506,0.00024170816],"domain_scores_gemma":[0.992945,0.0033407656,0.0007288776,0.0010130606,0.0017417844,0.00023044951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003722655,0.0013269451,0.0010349212,0.0044386615,0.00080189673,0.0035652218,0.0022987025,0.0014091388,0.0027453967],"category_scores_gemma":[0.015872004,0.00031880487,0.0015563426,0.0026451,0.00076900964,0.0040720673,0.0027281852,0.0020368844,0.0019539748],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033194345,0.00086583954,0.015585291,0.000661852,0.00016163883,0.0006123675,0.0011337297,0.07119151,0.011974731,0.016639564,0.009275997,0.8715656],"study_design_scores_gemma":[0.000043030752,0.00023037015,0.007652968,0.000109679546,0.00009100727,0.0002118589,0.0012549784,0.89835143,0.012351708,0.067233615,0.012399045,0.00007032127],"about_ca_topic_score_codex":0.007955666,"about_ca_topic_score_gemma":0.0102554625,"teacher_disagreement_score":0.007955666,"about_ca_system_score_codex":0.0017899035,"about_ca_system_score_gemma":0.0029591813,"threshold_uncertainty_score":0.019687533},"labels":[],"label_agreement":null},{"id":"W4285308692","doi":"10.18653/v1/2022.nlp4convai-1.17","title":"Toward Knowledge-Enriched Conversational Recommendation Systems","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Research (Canada)","funders":"National Research Foundation Singapore; National Research Foundation","keywords":"Computer science; Recommender system; Knowledge graph; Natural language processing; Scale (ratio); Artificial intelligence; Information retrieval; World Wide Web; Human–computer interaction","score_opus":0.060685804570036965,"score_gpt":0.2693948918523879,"score_spread":0.20870908728235094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285308692","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021713559,0.0009140832,0.9517721,0.0012943073,0.00016222348,0.00021596333,0.0009879934,0.01809047,0.0048492295],"genre_scores_gemma":[0.32640454,0.000723442,0.6542936,0.0011861556,0.00029571683,0.0004426409,0.0058204834,0.0008923248,0.009941026],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99681926,0.0014285822,0.00017002305,0.00090634736,0.00046453596,0.00021127328],"domain_scores_gemma":[0.99546045,0.002127986,0.00018693971,0.0010966957,0.0008938141,0.0002341355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036996964,0.0013772027,0.0011504962,0.0014571866,0.0011727113,0.0023752665,0.0032875144,0.0025657278,0.0038356944],"category_scores_gemma":[0.012401157,0.00088755705,0.0011507499,0.0011275902,0.00074517314,0.006728048,0.003965241,0.0029030265,0.004991476],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001041797,0.0010374958,0.003822843,0.0008949892,0.0003888677,0.0007567078,0.0021044104,0.13107558,0.03970167,0.035502728,0.055537727,0.72813517],"study_design_scores_gemma":[0.000040123727,0.00006747548,0.00032818769,0.000038453847,0.00008344124,0.00011124677,0.00028089483,0.93910867,0.009816544,0.032971796,0.017115965,0.00003706961],"about_ca_topic_score_codex":0.008060611,"about_ca_topic_score_gemma":0.013511878,"teacher_disagreement_score":0.008060611,"about_ca_system_score_codex":0.0011383769,"about_ca_system_score_gemma":0.0018556392,"threshold_uncertainty_score":0.01956606},"labels":[],"label_agreement":null},{"id":"W4285310604","doi":"10.18653/v1/2022.nlp4convai-1.5","title":"Data Augmentation for Intent Classification with Off-the-shelf Large Language Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Classifier (UML); Training set; Task (project management); Labeled data; Language model; Machine learning; Artificial intelligence; Scarcity; Data quality; Data mining; Metric (unit)","score_opus":0.11258400597139151,"score_gpt":0.31340503453173507,"score_spread":0.20082102856034356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285310604","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.065960266,0.0013381523,0.8993245,0.0007263013,0.00048328165,0.00054563244,0.0029419942,0.026339777,0.0023400972],"genre_scores_gemma":[0.44573924,0.00039585293,0.5334443,0.00079766946,0.00026602173,0.0015996513,0.012669299,0.0011087621,0.003979307],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99826956,0.0007417953,0.00011854482,0.0005066764,0.00025476122,0.000108696244],"domain_scores_gemma":[0.99315757,0.0043083206,0.00027214782,0.0012400239,0.0008153748,0.00020657193],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027702674,0.0019522941,0.0011862811,0.0012533248,0.0005886601,0.0011457518,0.0023866482,0.0015556375,0.0041594957],"category_scores_gemma":[0.015081982,0.000699079,0.0014994886,0.0010900898,0.0007981821,0.002879291,0.002618524,0.00395413,0.00504744],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010264708,0.00092985993,0.0048448523,0.0007414691,0.00015721584,0.00033146466,0.00083656947,0.05887706,0.04955791,0.00353752,0.023344345,0.8558152],"study_design_scores_gemma":[0.00012988893,0.00046449853,0.0018564285,0.00008563,0.0000701619,0.00018814622,0.0002857817,0.94735533,0.028079702,0.011374346,0.010035959,0.00007414851],"about_ca_topic_score_codex":0.0019418963,"about_ca_topic_score_gemma":0.0040321928,"teacher_disagreement_score":0.0041594957,"about_ca_system_score_codex":0.0006769065,"about_ca_system_score_gemma":0.0012616232,"threshold_uncertainty_score":0.0146507025},"labels":[],"label_agreement":null},{"id":"W4285464655","doi":"10.32920/ryerson.14660694.v1","title":"Identity matching in social media platforms","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; York University","funders":"","keywords":"Computer science; Matching (statistics); Social media; Identity (music); Set (abstract data type); Task (project management); World Wide Web; Information retrieval; Similarity (geometry); Process (computing); Social Semantic Web; String (physics); Semantic Web; Artificial intelligence; Engineering; Mathematics","score_opus":0.05234093860938426,"score_gpt":0.2957890022646055,"score_spread":0.24344806365522126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285464655","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1724203,0.001019554,0.80177623,0.0015882768,0.00017404693,0.0006201903,0.0012155365,0.0010252198,0.020160604],"genre_scores_gemma":[0.8355941,0.0006506011,0.15518384,0.00019824538,0.00013812925,0.0003368989,0.0014503396,0.00012307832,0.006324733],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99161536,0.003149416,0.0005086066,0.0016833744,0.0023724935,0.0006706715],"domain_scores_gemma":[0.99319446,0.0031212703,0.0012927157,0.0012744144,0.000878325,0.00023893663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059501356,0.00050111936,0.0007329164,0.006718991,0.0027531034,0.004633079,0.0016473087,0.0018694041,0.0027090993],"category_scores_gemma":[0.019349257,0.0005228519,0.0015358033,0.0058860807,0.0014828806,0.009702651,0.004405284,0.0010566682,0.0013176805],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043466882,0.0005115312,0.049613122,0.00050806865,0.00032848344,0.0018461078,0.005694241,0.0498074,0.008649423,0.5280575,0.009806506,0.34474292],"study_design_scores_gemma":[0.000026060325,0.000117647796,0.015539788,0.0002235876,0.00013297923,0.001198767,0.004829665,0.51927143,0.012109199,0.40417832,0.04226056,0.000111958594],"about_ca_topic_score_codex":0.005532076,"about_ca_topic_score_gemma":0.0033715777,"teacher_disagreement_score":0.006718991,"about_ca_system_score_codex":0.0019169014,"about_ca_system_score_gemma":0.001636804,"threshold_uncertainty_score":0.031467736},"labels":[],"label_agreement":null},{"id":"W4285817825","doi":"10.1007/978-3-031-10986-7_17","title":"CLINER: Clinical Interrogation Named Entity Recognition","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Interrogation; Dialog box; Computer science; Named-entity recognition; Information extraction; Context (archaeology); Task (project management); Artificial intelligence; Exploit; Natural language processing; Information retrieval; World Wide Web; Computer security","score_opus":0.06366758333984053,"score_gpt":0.3129572673007279,"score_spread":0.2492896839608874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285817825","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006864625,0.0034854873,0.6776112,0.0028776808,0.0018556347,0.00078355067,0.036475755,0.19989914,0.07014699],"genre_scores_gemma":[0.07278485,0.002580227,0.7043473,0.0027837579,0.001108426,0.0009188334,0.093493804,0.01849214,0.10349066],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99876785,0.00029729784,0.00013425616,0.00032105733,0.0004038152,0.000075625685],"domain_scores_gemma":[0.9982584,0.00092293636,0.00008777102,0.00036813057,0.00025786663,0.0001048055],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021542902,0.0013696287,0.0009169831,0.0022702096,0.00043331087,0.0025090193,0.002001497,0.0013619869,0.058210466],"category_scores_gemma":[0.0046618357,0.00069073006,0.0006683815,0.0019043674,0.0004978205,0.0029506919,0.0023475394,0.0015645467,0.050774958],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017867192,0.00006321251,0.0006785403,0.00038239264,0.00004549185,0.0003249512,0.00017739023,0.001281675,0.0068223793,0.010496663,0.44218466,0.53736407],"study_design_scores_gemma":[0.00010109279,0.00011611642,0.0022192637,0.00030158577,0.00008765763,0.002673708,0.00020232379,0.054896567,0.04079043,0.042623684,0.8558535,0.00013398853],"about_ca_topic_score_codex":0.0010524923,"about_ca_topic_score_gemma":0.0014105836,"teacher_disagreement_score":0.058210466,"about_ca_system_score_codex":0.00052128267,"about_ca_system_score_gemma":0.0009021333,"threshold_uncertainty_score":0.19473338},"labels":[],"label_agreement":null},{"id":"W4285891774","doi":"10.1007/s10489-022-03944-z","title":"Reward Modeling for Mitigating Toxicity in Transformer-based Language Models","year":2022,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Language model; Transformer; Artificial intelligence; Machine learning; Natural language processing; Engineering","score_opus":0.08285491862021747,"score_gpt":0.1938981566713769,"score_spread":0.11104323805115941,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285891774","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044246536,0.00019077865,0.9523178,0.00044050993,0.000042497933,0.00009572236,0.000117974974,0.0014111712,0.0011371908],"genre_scores_gemma":[0.8419373,0.0001974949,0.15306467,0.00041264662,0.000051988154,0.00025761456,0.00037311277,0.00023864588,0.0034664855],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881077,0.00063130923,0.00006515681,0.00023283804,0.00017475439,0.00008512456],"domain_scores_gemma":[0.9944125,0.004335142,0.00033732058,0.000284047,0.0004812772,0.00014964357],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029024126,0.00083294255,0.00069377007,0.00052941515,0.0003732193,0.00076302607,0.0012106397,0.0008079373,0.0020737615],"category_scores_gemma":[0.011790114,0.00037865285,0.0006762123,0.00034918566,0.0007439414,0.0020205202,0.0011927396,0.0019491331,0.0006879164],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046313362,0.00020650723,0.0034289183,0.00020388076,0.00007409819,0.0002129024,0.0002994216,0.7885953,0.010704401,0.02969086,0.0031151078,0.16300549],"study_design_scores_gemma":[0.000013700752,0.000034476485,0.00007583073,0.0000032157247,0.000008218088,0.00001578113,0.0000065713543,0.9922069,0.0012336994,0.006093194,0.00030343691,0.000004941536],"about_ca_topic_score_codex":0.003141577,"about_ca_topic_score_gemma":0.00505238,"teacher_disagreement_score":0.003141577,"about_ca_system_score_codex":0.001032476,"about_ca_system_score_gemma":0.0011617906,"threshold_uncertainty_score":0.015349627},"labels":[],"label_agreement":null},{"id":"W4286219887","doi":"10.1002/essoar.10510978.2","title":"m-NLP inference models using simulation and regression techniques","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Norges Forskningsråd; H2020 European Research Council; China Scholarship Council; Compute Canada","keywords":"Preprint; Inference; Computer science; World Wide Web; Information retrieval; Natural language processing; Artificial intelligence","score_opus":0.1589566516328696,"score_gpt":0.38556881623716716,"score_spread":0.22661216460429756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286219887","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0089289965,0.0008008593,0.9831692,0.0013114284,0.00014465174,0.00013748089,0.001352855,0.0018344392,0.002320104],"genre_scores_gemma":[0.32169658,0.0015848189,0.6514734,0.0008379227,0.00065623136,0.001523721,0.008469348,0.0015259773,0.012232011],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9911987,0.0064381915,0.00030326133,0.0011727037,0.0005906453,0.00029637135],"domain_scores_gemma":[0.927499,0.064970136,0.0015081527,0.0030873595,0.0024365212,0.00049881754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013584917,0.0015819195,0.00232077,0.0023439093,0.0013126656,0.0032857787,0.0037861464,0.002740642,0.010607422],"category_scores_gemma":[0.072668605,0.0016110421,0.0029019434,0.0030241536,0.001473301,0.0040516974,0.0025244541,0.00516212,0.0037293986],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031569856,0.00015039986,0.0048215794,0.00032996316,0.00046720137,0.00021724147,0.00022355594,0.79216784,0.00027915326,0.09672404,0.016075838,0.0882275],"study_design_scores_gemma":[0.00003291301,0.00001006981,0.00016550577,0.000023803628,0.000021462109,0.000013184081,0.000015178518,0.9552113,0.0001159609,0.04299266,0.0013878042,0.000010249145],"about_ca_topic_score_codex":0.029265175,"about_ca_topic_score_gemma":0.028014107,"teacher_disagreement_score":0.029265175,"about_ca_system_score_codex":0.0020918567,"about_ca_system_score_gemma":0.0036832043,"threshold_uncertainty_score":0.07184476},"labels":[],"label_agreement":null},{"id":"W4286531984","doi":"10.1109/saner53432.2022.00039","title":"Evaluating the Use of Semantics for Identifying Task-relevant Textual Information","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Software Analysis, Evolution and Reengineering (SANER)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Semantics (computer science); Task (project management); Information retrieval; Natural language processing; World Wide Web; Artificial intelligence; Programming language; Engineering","score_opus":0.12738109810621479,"score_gpt":0.3389998145211017,"score_spread":0.21161871641488694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286531984","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.860563,0.004505025,0.11574783,0.000668478,0.00023944565,0.00075805763,0.002749993,0.0068310965,0.007937041],"genre_scores_gemma":[0.8435463,0.00087347045,0.1486559,0.00016262196,0.00011668778,0.00030424836,0.004414141,0.0002850117,0.0016416813],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99019194,0.0040553724,0.0013490586,0.0015027205,0.0025221612,0.00037876912],"domain_scores_gemma":[0.8819098,0.09993708,0.005724637,0.003512865,0.0077513377,0.0011642107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010239484,0.0018464132,0.00078252814,0.012310308,0.0007135638,0.0024310274,0.0012118183,0.0022525545,0.00085828535],"category_scores_gemma":[0.06738358,0.00032443678,0.0011375849,0.0033447538,0.0009902626,0.005340454,0.0016051166,0.001150603,0.00068452617],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040093446,0.0026091342,0.15822554,0.004959686,0.0010778093,0.00078657316,0.0063662957,0.027188847,0.052640922,0.004318273,0.0098090265,0.72800857],"study_design_scores_gemma":[0.00076833257,0.0061597023,0.14128886,0.0009701589,0.0019056408,0.0026061335,0.009224363,0.6784028,0.113276094,0.01546946,0.029489886,0.00043864903],"about_ca_topic_score_codex":0.0044277324,"about_ca_topic_score_gemma":0.0064677657,"teacher_disagreement_score":0.012310308,"about_ca_system_score_codex":0.00085336197,"about_ca_system_score_gemma":0.0015234287,"threshold_uncertainty_score":0.05415219},"labels":[],"label_agreement":null},{"id":"W4286903770","doi":"10.48550/arxiv.2110.07752","title":"Hindsight: Posterior-guided training of retrievers for improved\\n open-ended generation","year":2021,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Hindsight bias; Computer science; Context (archaeology); Generator (circuit theory); Relevance (law); Labrador Retriever; Artificial intelligence; Natural language processing; Information retrieval; Psychology; Cognitive psychology; Medicine; Power (physics); Geography","score_opus":0.25625230964833545,"score_gpt":0.23918258047907628,"score_spread":0.017069729169259168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286903770","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09677041,0.002081512,0.829077,0.0009902449,0.0004797649,0.0004540634,0.0011020685,0.05854333,0.010501524],"genre_scores_gemma":[0.6117968,0.0003322705,0.3536959,0.0021696314,0.00025226016,0.00072064623,0.0067073414,0.003299398,0.021025762],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998473,0.0005185162,0.00008463687,0.0005723888,0.00019562613,0.00015582106],"domain_scores_gemma":[0.9952887,0.0033410701,0.00017049292,0.0006155938,0.00040896083,0.00017510072],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033723195,0.0023113436,0.0012391402,0.0007408451,0.00061182585,0.0013362837,0.003588617,0.0028860755,0.012838306],"category_scores_gemma":[0.010256765,0.00095043384,0.0011691321,0.00046795636,0.0012984928,0.0023919584,0.002977049,0.0037473533,0.0054806545],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015347484,0.00067339826,0.004503513,0.000655142,0.00025142808,0.0006684771,0.0007080434,0.35088044,0.018955002,0.00656953,0.03193069,0.5826696],"study_design_scores_gemma":[0.00015709572,0.0002458648,0.00043708587,0.000040492032,0.000040046096,0.00014796644,0.000063674786,0.980164,0.010462147,0.0044462346,0.0037682287,0.000027130518],"about_ca_topic_score_codex":0.0045990143,"about_ca_topic_score_gemma":0.008635299,"teacher_disagreement_score":0.012838306,"about_ca_system_score_codex":0.000784857,"about_ca_system_score_gemma":0.0012552005,"threshold_uncertainty_score":0.042948425},"labels":[],"label_agreement":null},{"id":"W4287120901","doi":"10.1162/tacl_a_00492","title":"Generate, Annotate, and Learn: NLP with Synthetic Text","year":2022,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada)","funders":"Defense Advanced Research Projects Agency","keywords":"Computer science; Artificial intelligence; Transformer; Natural language processing; Classifier (UML); Machine learning; Labeled data; Task (project management); Distillation; Language model","score_opus":0.012829661286926128,"score_gpt":0.2278401156433895,"score_spread":0.21501045435646338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287120901","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17636134,0.0007228936,0.7923977,0.0022512237,0.00040315287,0.00028205922,0.003608165,0.013176553,0.010796974],"genre_scores_gemma":[0.6470042,0.00016514186,0.3397616,0.00041239272,0.000099164674,0.00032401734,0.00590084,0.0007953396,0.0055372184],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982247,0.0010025067,0.000056768735,0.0003669994,0.00027577847,0.00007326284],"domain_scores_gemma":[0.99004793,0.007523117,0.00024122754,0.0013813397,0.00063211535,0.00017425862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027076362,0.0008602803,0.00048027412,0.00053639576,0.00056359987,0.0010477084,0.0014526291,0.0012586011,0.0040445435],"category_scores_gemma":[0.015430701,0.0003060373,0.0005009374,0.0005995124,0.0012895856,0.0031716696,0.0019122087,0.0018574374,0.0016790696],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010519257,0.0006229797,0.003493272,0.0006655763,0.00013124397,0.00054170965,0.0008018955,0.57675785,0.012920911,0.063366696,0.0351302,0.3045157],"study_design_scores_gemma":[0.000050557283,0.000051558174,0.00015482171,0.000016807686,0.000006946076,0.000040284183,0.00006872984,0.9605917,0.0093613975,0.025367115,0.004277063,0.000012997452],"about_ca_topic_score_codex":0.0021027985,"about_ca_topic_score_gemma":0.0036578614,"teacher_disagreement_score":0.0040445435,"about_ca_system_score_codex":0.0008224361,"about_ca_system_score_gemma":0.0006645097,"threshold_uncertainty_score":0.014319479},"labels":[],"label_agreement":null},{"id":"W4287122359","doi":"10.48550/arxiv.2106.05346","title":"End-to-End Training of Multi-Document Reader and Retriever for\\n Open-Domain Question Answering","year":2021,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Question answering; Benchmark (surveying); Information retrieval; Domain (mathematical analysis); Set (abstract data type); Training set; Open domain; Artificial intelligence; Labrador Retriever; Relevance (law); Machine learning; Mathematics","score_opus":0.14848206269298694,"score_gpt":0.24423374413988705,"score_spread":0.0957516814469001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287122359","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018458292,0.0006531227,0.96105576,0.00034489698,0.0000897112,0.00021666108,0.000374613,0.016774157,0.0020327338],"genre_scores_gemma":[0.23556611,0.00033145017,0.74317056,0.0008318265,0.00019246359,0.0005735313,0.004147909,0.001012949,0.014173231],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984072,0.00044318053,0.00009636638,0.00066932046,0.00023899977,0.00014491161],"domain_scores_gemma":[0.9970692,0.0016057102,0.00013842627,0.00054955913,0.0004976548,0.00013940997],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024742316,0.0019685163,0.0016722447,0.0010238074,0.00067720003,0.0013745982,0.0039326847,0.0036793614,0.010847271],"category_scores_gemma":[0.007251043,0.0007951137,0.0013188098,0.0008066526,0.00082015555,0.0032988286,0.0025709074,0.0032385758,0.008533395],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005950548,0.0005593201,0.001719082,0.0004465101,0.00014306097,0.00030520122,0.00047569824,0.09038835,0.024764573,0.0042371443,0.018676259,0.8576897],"study_design_scores_gemma":[0.00006165457,0.00018928801,0.00044022448,0.000020635023,0.000044952038,0.0001685487,0.00008408885,0.97482306,0.014397101,0.004674574,0.0050732726,0.000022601773],"about_ca_topic_score_codex":0.0038233974,"about_ca_topic_score_gemma":0.008648984,"teacher_disagreement_score":0.010847271,"about_ca_system_score_codex":0.0010561595,"about_ca_system_score_gemma":0.0013469998,"threshold_uncertainty_score":0.036287725},"labels":[],"label_agreement":null},{"id":"W4287240206","doi":"10.48550/arxiv.2104.01940","title":"What's the best place for an AI conference, Vancouver or ______: Why\\n completing comparative questions is difficult","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Competence (human resources); Artificial intelligence; sort; Natural language processing; Set (abstract data type); Benchmark (surveying); Task (project management); Language model; Machine learning; Cognitive science; Information retrieval; Psychology","score_opus":0.21797182147043911,"score_gpt":0.2610461724021927,"score_spread":0.043074350931753574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287240206","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034781344,0.011466859,0.007539747,0.11427778,0.008191057,0.00019906506,0.0051290835,0.002080483,0.8163346],"genre_scores_gemma":[0.26795354,0.010616132,0.020205298,0.008594787,0.0014864148,0.0001672401,0.008091654,0.0010045866,0.68188035],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99937195,0.00016752392,0.000029848066,0.00014551482,0.00015480918,0.00013033419],"domain_scores_gemma":[0.99819046,0.0003140458,0.000057194015,0.000119555494,0.0006777131,0.000641022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012137764,0.0006275238,0.00045993965,0.00090904644,0.0052022687,0.005831003,0.00088694924,0.001947135,0.10909859],"category_scores_gemma":[0.005072177,0.00031298082,0.00030966874,0.0015480467,0.0012575034,0.0037458385,0.0011559047,0.0025699623,0.032306895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018435069,0.00009560121,0.0045057526,0.00041641525,0.000042110587,0.0003261049,0.0014839144,0.000918353,0.0018516252,0.015962003,0.7516475,0.22256626],"study_design_scores_gemma":[0.000023885736,0.000036302838,0.008327434,0.0003508819,0.000022804476,0.00018930715,0.004559695,0.0012607544,0.0013632488,0.008251397,0.97556245,0.000051784555],"about_ca_topic_score_codex":0.20456882,"about_ca_topic_score_gemma":0.5441618,"teacher_disagreement_score":0.20456882,"about_ca_system_score_codex":0.004923104,"about_ca_system_score_gemma":0.005168947,"threshold_uncertainty_score":0.40675616},"labels":[],"label_agreement":null},{"id":"W4287551482","doi":"10.48550/arxiv.2012.09936","title":"Named Entity Recognition in the Legal Domain using a Pointer Generator\\n Network","year":2020,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Named-entity recognition; Computer science; Pointer (user interface); Security token; Natural language processing; Task (project management); Artificial intelligence; Annotation; Information retrieval; Computer security","score_opus":0.13084094898160611,"score_gpt":0.19791716279038815,"score_spread":0.06707621380878204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287551482","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11241019,0.0009909734,0.865303,0.0014778618,0.00024304564,0.00019637495,0.0015314085,0.009477593,0.008369644],"genre_scores_gemma":[0.6703828,0.0008508585,0.30204034,0.00063910463,0.00021986243,0.000310673,0.006379358,0.0003301796,0.018846758],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9996716,0.0000733854,0.000016623606,0.00015292862,0.00004843823,0.00003716077],"domain_scores_gemma":[0.99915147,0.00048264474,0.000056949582,0.00012930509,0.00014547564,0.000034171946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000992773,0.00058479037,0.0004362096,0.0011055103,0.000546935,0.00085626164,0.0014192442,0.0013006576,0.0032772755],"category_scores_gemma":[0.0030959963,0.00034979056,0.00064727885,0.0010679751,0.00057278917,0.0026832535,0.001184803,0.0013052772,0.0017214736],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003235228,0.00033741616,0.0032014681,0.00016110577,0.000120473975,0.00049956795,0.00026019028,0.384127,0.015701182,0.016880691,0.019320462,0.55906695],"study_design_scores_gemma":[0.000007547201,0.00002711082,0.00041175485,0.000009551782,0.000014733588,0.000041792864,0.000018469873,0.9892217,0.003434802,0.005203042,0.0016005891,0.000009020301],"about_ca_topic_score_codex":0.009096192,"about_ca_topic_score_gemma":0.01279153,"teacher_disagreement_score":0.009096192,"about_ca_system_score_codex":0.0009540298,"about_ca_system_score_gemma":0.0007524603,"threshold_uncertainty_score":0.018086493},"labels":[],"label_agreement":null},{"id":"W4287854991","doi":"10.18653/v1/2022.naacl-main.170","title":"Towards Understanding Large-Scale Discourse Structures in Pre-Trained and Fine-Tuned Language Models","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Redundancy (engineering); Natural language processing; Artificial intelligence; Language model; Scale (ratio); Simple (philosophy)","score_opus":0.023288364856999623,"score_gpt":0.2689538667565817,"score_spread":0.2456655018995821,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287854991","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2608354,0.0018110676,0.7304624,0.0013253568,0.00009816852,0.00013232513,0.0007424067,0.0022472467,0.0023455604],"genre_scores_gemma":[0.80141044,0.00043786367,0.19250463,0.00031782355,0.00014598558,0.00022174141,0.0022821338,0.00042847116,0.0022509452],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99705577,0.0017302402,0.00010907197,0.0007220481,0.00022287272,0.00015996513],"domain_scores_gemma":[0.9694751,0.025603611,0.0009741306,0.0024626392,0.001105324,0.0003791635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006991529,0.0013601953,0.0011326056,0.0019055171,0.00081653654,0.0033979896,0.0019138204,0.0022674948,0.0012023723],"category_scores_gemma":[0.033731423,0.0012758012,0.0008792107,0.0011367459,0.0011847008,0.0069360957,0.0020413694,0.0058553377,0.0007326081],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006259146,0.00035736142,0.008395908,0.00029328378,0.0003951127,0.00017910791,0.0024078977,0.7946162,0.015147035,0.014193109,0.0028092,0.1605798],"study_design_scores_gemma":[0.00001739682,0.00003991053,0.0006486861,0.000018322611,0.000016371534,0.000012584758,0.00012714374,0.9854036,0.0014741248,0.011599534,0.0006265795,0.000015693773],"about_ca_topic_score_codex":0.009543543,"about_ca_topic_score_gemma":0.018751027,"teacher_disagreement_score":0.009543543,"about_ca_system_score_codex":0.0025178778,"about_ca_system_score_gemma":0.0018184502,"threshold_uncertainty_score":0.036975205},"labels":[],"label_agreement":null},{"id":"W4287887656","doi":"10.18653/v1/2022.naacl-main.114","title":"Same Neurons, Different Languages: Probing Morphosyntax in Multilingual Pre-trained Models","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Vetenskapsrådet; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; National Science Foundation","keywords":"Computer science; Linguistics; Artificial intelligence; Natural language processing; Philosophy","score_opus":0.01915628939806864,"score_gpt":0.25428684507598714,"score_spread":0.23513055567791852,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287887656","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7409058,0.0038033116,0.21351795,0.0025973225,0.0007781226,0.00012490537,0.0056330096,0.0088521615,0.023787478],"genre_scores_gemma":[0.95760083,0.0003787292,0.029971365,0.00026797576,0.000069090835,0.00008205803,0.005720489,0.0017897426,0.0041197445],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986669,0.00049005164,0.00007454786,0.00047459584,0.00012693553,0.00016709199],"domain_scores_gemma":[0.9939851,0.0044835475,0.000119961514,0.0007253057,0.0004916982,0.00019434498],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020718626,0.0013137131,0.0007623113,0.000863812,0.0008403775,0.0035550452,0.002071627,0.001535385,0.006097706],"category_scores_gemma":[0.010817234,0.001132556,0.0011672082,0.0012326527,0.0010437538,0.0057517462,0.003519366,0.00442494,0.002941853],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035526103,0.0006532403,0.0508715,0.00079729356,0.0012428603,0.0015380566,0.0061215004,0.15636572,0.038305074,0.020013735,0.032291967,0.6882464],"study_design_scores_gemma":[0.00020299555,0.00023678098,0.01521495,0.00019980418,0.00046963053,0.0007737657,0.0027762137,0.88172257,0.026035333,0.059286725,0.0129149025,0.00016639316],"about_ca_topic_score_codex":0.012678276,"about_ca_topic_score_gemma":0.03495999,"teacher_disagreement_score":0.012678276,"about_ca_system_score_codex":0.0014304713,"about_ca_system_score_gemma":0.0011986678,"threshold_uncertainty_score":0.02520895},"labels":[],"label_agreement":null},{"id":"W4287889735","doi":"10.18653/v1/2022.findings-naacl.151","title":"Great Power, Great Responsibility: Recommendations for Reducing Energy for Training Language Models","year":2022,"lang":"en","type":"article","venue":"Findings of the Association for Computational Linguistics: NAACL 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Massachusetts Green High Performance Computing Center; Leukemia and Lymphoma Society of Canada; National Science Foundation","keywords":"Computer science; Energy consumption; Inference; Pace; Cloud computing; Transformer; Efficient energy use; Machine learning; Artificial intelligence; Risk analysis (engineering); Engineering","score_opus":0.03769841022498763,"score_gpt":0.29149186124092274,"score_spread":0.2537934510159351,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287889735","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044264193,0.045963854,0.2954997,0.576948,0.0075787874,0.00056601677,0.002152587,0.013649608,0.053215016],"genre_scores_gemma":[0.08480859,0.06035386,0.67628574,0.10067007,0.004559816,0.0022521634,0.004844624,0.007806075,0.05841904],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99483323,0.0023193494,0.00035415197,0.00052261807,0.0015862206,0.0003844307],"domain_scores_gemma":[0.96031874,0.02444462,0.00082213397,0.003493205,0.0083966255,0.0025246877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00952204,0.0020443585,0.0011778746,0.0029080098,0.0018693338,0.0058419546,0.00773802,0.006517267,0.0493919],"category_scores_gemma":[0.05165201,0.001655877,0.0013954372,0.0043852483,0.0033481666,0.020110983,0.003375087,0.009638011,0.024651064],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018452817,0.0003196915,0.0011321657,0.0010314243,0.000064507316,0.000108614535,0.00029755654,0.0071336864,0.0015853194,0.05173662,0.60894597,0.32746005],"study_design_scores_gemma":[0.0003150504,0.00012334644,0.0014413437,0.0016347109,0.000086418535,0.00023168103,0.0010205891,0.028409865,0.0029720473,0.15742996,0.8061858,0.00014924223],"about_ca_topic_score_codex":0.01761981,"about_ca_topic_score_gemma":0.03374566,"teacher_disagreement_score":0.0493919,"about_ca_system_score_codex":0.002918367,"about_ca_system_score_gemma":0.006886862,"threshold_uncertainty_score":0.16523236},"labels":[],"label_agreement":null},{"id":"W4287890446","doi":"10.18653/v1/2022.naacl-industry.30","title":"Adversarial Text Normalization","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Adversarial system; Normalization (sociology); Computer science; Computational linguistics; Artificial intelligence; Maya; Natural language processing; Linguistics; History; Sociology; Philosophy; Anthropology; Archaeology","score_opus":0.013751583047096661,"score_gpt":0.21579982539512862,"score_spread":0.20204824234803195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287890446","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0055060145,0.0037572403,0.9595456,0.0027435417,0.0028116363,0.00020745085,0.0025280546,0.009373717,0.013526671],"genre_scores_gemma":[0.27968058,0.005504437,0.5407698,0.00415281,0.0038646813,0.00096086285,0.022068359,0.004669709,0.13832878],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99784195,0.0007507367,0.000079410864,0.0006727565,0.00048525227,0.00016986846],"domain_scores_gemma":[0.9969543,0.001414261,0.000117881544,0.0009724439,0.0004435433,0.00009757394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028483788,0.0018458967,0.0013322535,0.0012708161,0.00083178736,0.0018555649,0.0022872994,0.001985884,0.017032236],"category_scores_gemma":[0.009871025,0.00073207595,0.0011651387,0.0014296112,0.00108575,0.0027367754,0.003483934,0.0033931115,0.022756154],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038321625,0.00018146829,0.0005733063,0.0003467645,0.00021525123,0.0002836719,0.00012196981,0.09833627,0.014568323,0.035700683,0.28180555,0.5674835],"study_design_scores_gemma":[0.00004487464,0.00007506135,0.00041304025,0.000088160596,0.00005135682,0.00033655373,0.000059787933,0.864791,0.015108767,0.059548147,0.05943815,0.00004515583],"about_ca_topic_score_codex":0.0022332827,"about_ca_topic_score_gemma":0.0033690168,"teacher_disagreement_score":0.017032236,"about_ca_system_score_codex":0.0006621274,"about_ca_system_score_gemma":0.0008856129,"threshold_uncertainty_score":0.056978524},"labels":[],"label_agreement":null},{"id":"W4287891037","doi":"10.18653/v1/2022.naacl-main.188","title":"How Gender Debiasing Affects Internal Model Representations, and Why It Matters","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation","keywords":"Debiasing; Computer science; Data science; Linguistics; Cognitive science; Psychology; Philosophy","score_opus":0.026778599557794835,"score_gpt":0.25861221761684694,"score_spread":0.23183361805905212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287891037","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6859344,0.0030314757,0.17580047,0.02112909,0.0011865057,0.00013050006,0.0033304854,0.0047760145,0.104681045],"genre_scores_gemma":[0.9789109,0.00039444817,0.012659568,0.00063867704,0.00008514809,0.000024664405,0.0011911559,0.0014567982,0.0046386225],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99748516,0.0011000105,0.000107168526,0.00061288656,0.0004651199,0.00022958755],"domain_scores_gemma":[0.98122257,0.010022571,0.0011322794,0.003975575,0.0032036016,0.00044346048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004941644,0.00049154495,0.00041870464,0.000750043,0.0010013259,0.006259785,0.0010766663,0.0014094972,0.015387665],"category_scores_gemma":[0.077773504,0.0005870132,0.0005270726,0.0007897764,0.0016027777,0.012527985,0.00205473,0.0019576822,0.005289509],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026902906,0.00036114224,0.14588134,0.000498537,0.00033815607,0.0010115714,0.023422975,0.021303121,0.03443258,0.14732778,0.07095989,0.5517726],"study_design_scores_gemma":[0.0002781969,0.0004886402,0.084352195,0.0006534795,0.00073529314,0.0019912503,0.027449302,0.18638995,0.03893632,0.58960754,0.068715826,0.00040210588],"about_ca_topic_score_codex":0.008081241,"about_ca_topic_score_gemma":0.006203726,"teacher_disagreement_score":0.015387665,"about_ca_system_score_codex":0.0011103061,"about_ca_system_score_gemma":0.0008299688,"threshold_uncertainty_score":0.051476836},"labels":[],"label_agreement":null},{"id":"W4288072950","doi":"10.1145/3459637.3482352","title":"Identifying Untrustworthy Samples: Data Filtering for Open-domain Dialogues with Bayesian Optimization","year":2021,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science","score_opus":0.2312629661684371,"score_gpt":0.2366491494227362,"score_spread":0.005386183254299098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4288072950","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03107569,0.001119162,0.9616073,0.0006181667,0.00007821852,0.00019501656,0.00035254017,0.004206547,0.0007474027],"genre_scores_gemma":[0.30037543,0.0002962032,0.6921962,0.0008759492,0.00012316216,0.0005626025,0.002854386,0.00063010596,0.0020859807],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9924908,0.0040254164,0.00043062476,0.001757802,0.000954008,0.00034129136],"domain_scores_gemma":[0.9844305,0.011186991,0.0007330607,0.0015120265,0.001641737,0.0004956325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009207808,0.0023198095,0.002471533,0.001732195,0.001371368,0.0018739963,0.0030640576,0.0027665286,0.0013888102],"category_scores_gemma":[0.02594616,0.001160397,0.002068285,0.0010347901,0.0012449007,0.0030833962,0.0027250163,0.004579747,0.0013725635],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014766322,0.001123257,0.010030111,0.0006814968,0.00068987923,0.000279486,0.002007018,0.26208594,0.021006824,0.005284953,0.013852176,0.68148226],"study_design_scores_gemma":[0.00003885172,0.00006090733,0.00095300947,0.000028581095,0.000036211677,0.000047790847,0.00013189424,0.9886662,0.0036780213,0.0048198653,0.0015094576,0.000029149976],"about_ca_topic_score_codex":0.008483911,"about_ca_topic_score_gemma":0.011897356,"teacher_disagreement_score":0.009207808,"about_ca_system_score_codex":0.0013292012,"about_ca_system_score_gemma":0.002019874,"threshold_uncertainty_score":0.04869616},"labels":[],"label_agreement":null},{"id":"W4289523162","doi":"10.1073/pnas.2123433119","title":"A neural network solves, explains, and generates university math problems by program synthesis and few-shot learning at human level","year":2022,"lang":"en","type":"article","venue":"Proceedings of the National Academy of Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":182,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial neural network; Benchmark (surveying); Linear algebra; Code (set theory); Artificial intelligence; Algorithm; Algebra over a field; Machine learning; Theoretical computer science; Mathematics; Programming language; Pure mathematics; Set (abstract data type)","score_opus":0.07421882056051196,"score_gpt":0.269741645150327,"score_spread":0.19552282458981504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289523162","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33807468,0.0039125164,0.51808333,0.0031581996,0.0010429814,0.0008935091,0.009730894,0.096142255,0.028961666],"genre_scores_gemma":[0.61652344,0.000659593,0.33931208,0.0021974673,0.00014204635,0.000716312,0.025143575,0.0014435345,0.013861986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914134,0.00015811234,0.000037784783,0.00046824,0.00012391446,0.00007065983],"domain_scores_gemma":[0.99797815,0.0012846907,0.000083546714,0.0003043874,0.00025857647,0.00009078075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010121053,0.0020063873,0.0005357964,0.00085615274,0.00042100184,0.0010979903,0.0023295782,0.0019058399,0.006858732],"category_scores_gemma":[0.006624418,0.00049612933,0.0012530085,0.00059224974,0.0005954323,0.001848836,0.00097422814,0.0024726912,0.0034059775],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005678234,0.0008405251,0.0085090585,0.0010846345,0.0002968543,0.0003444401,0.00042252886,0.18562621,0.021265747,0.004707247,0.052587688,0.7237472],"study_design_scores_gemma":[0.000073696654,0.00016494309,0.0013230451,0.000048698297,0.000046945337,0.000088223955,0.00008222945,0.97527254,0.010469045,0.005634849,0.006773788,0.000021816457],"about_ca_topic_score_codex":0.011478829,"about_ca_topic_score_gemma":0.022569984,"teacher_disagreement_score":0.011478829,"about_ca_system_score_codex":0.0016385823,"about_ca_system_score_gemma":0.0015722115,"threshold_uncertainty_score":0.022944748},"labels":[],"label_agreement":null},{"id":"W4289593661","doi":"10.1093/biosci/biac062","title":"Overcoming Language Barriers in Academia: Machine Translation Tools and a Vision for a Multilingual Future","year":2022,"lang":"en","type":"article","venue":"BioScience","topic":"Topic Modeling","field":"Computer Science","cited_by":137,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"University of Ottawa","keywords":"Machine translation; Computer science; Perspective (graphical); Translation (biology); Data science; Knowledge management; Engineering ethics; Artificial intelligence; Engineering","score_opus":0.021741262860367767,"score_gpt":0.31281036398085654,"score_spread":0.2910691011204888,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289593661","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017561018,0.034750488,0.45566982,0.41020077,0.004369293,0.00021238894,0.0006067519,0.0029829224,0.07364653],"genre_scores_gemma":[0.38025585,0.028580641,0.53532815,0.024616167,0.0042606727,0.0012303401,0.0018608347,0.0016602473,0.022207148],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9816434,0.013233941,0.0010462011,0.0014159961,0.00184465,0.0008159475],"domain_scores_gemma":[0.9430419,0.036061805,0.0028296218,0.0069035944,0.0078329975,0.003330061],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.045296263,0.001109195,0.001342218,0.0063004256,0.006050266,0.020456953,0.0033482285,0.0063763615,0.010084802],"category_scores_gemma":[0.060118,0.00089384627,0.001251877,0.0063667735,0.01193702,0.041704323,0.0106474515,0.007911871,0.0056526964],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009312291,0.00010564956,0.0022082303,0.0010412431,0.000054819153,0.00035759842,0.011509551,0.0023125522,0.0010766876,0.73907804,0.03470759,0.2074549],"study_design_scores_gemma":[0.000027602096,0.000033988737,0.00053907034,0.001205629,0.000035192083,0.00022697494,0.0052055065,0.0074922414,0.00091904163,0.7927182,0.19153072,0.00006578265],"about_ca_topic_score_codex":0.004088724,"about_ca_topic_score_gemma":0.0043074237,"teacher_disagreement_score":0.97954303,"about_ca_system_score_codex":0.006196445,"about_ca_system_score_gemma":0.014497229,"threshold_uncertainty_score":0.23955244},"labels":[],"label_agreement":null},{"id":"W4289670471","doi":"10.48550/arxiv.1809.01494","title":"Interpretation of Natural Language Rules in Conversational Machine\\n Reading","year":2018,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Task (project management); Reading (process); Interpretation (philosophy); Question answering; Natural (archaeology); Artificial intelligence; Government (linguistics); Natural language processing; Linguistics","score_opus":0.036903497926397306,"score_gpt":0.2000435815320063,"score_spread":0.163140083605609,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289670471","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14169207,0.004841885,0.8003364,0.004067578,0.0006579473,0.0013256911,0.004166761,0.01877347,0.024138136],"genre_scores_gemma":[0.6586084,0.0009896441,0.32511622,0.0011845494,0.0003625577,0.0006562818,0.008669106,0.00076097355,0.0036522064],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98199075,0.0112607945,0.00087476254,0.003431348,0.0018144379,0.00062777306],"domain_scores_gemma":[0.9568505,0.03417097,0.0014139423,0.0039383657,0.003042345,0.00058383803],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010248698,0.0019251644,0.0015015691,0.00348558,0.001946293,0.0066044126,0.0035010488,0.0030684872,0.0049983934],"category_scores_gemma":[0.046989225,0.0012406488,0.0020561698,0.0015324216,0.0028033627,0.0095652435,0.003612083,0.0039060423,0.0031577444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012629644,0.001033365,0.011402373,0.002526071,0.0005740944,0.001714155,0.016091907,0.13594323,0.021077622,0.049666606,0.027803011,0.73090464],"study_design_scores_gemma":[0.000097239514,0.00016991592,0.003581229,0.00020581229,0.00013108827,0.00030737734,0.002387724,0.79797804,0.015844034,0.16006361,0.019063033,0.0001709304],"about_ca_topic_score_codex":0.013493363,"about_ca_topic_score_gemma":0.011194825,"teacher_disagreement_score":0.013493363,"about_ca_system_score_codex":0.0022621427,"about_ca_system_score_gemma":0.002085229,"threshold_uncertainty_score":0.054201007},"labels":[],"label_agreement":null},{"id":"W4289711940","doi":"10.7717/peerj-cs.1066","title":"Causal graph extraction from news: a comparative study of time-series causality learning techniques","year":2022,"lang":"en","type":"article","venue":"PeerJ Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Agencia Nacional de Promoción Científica y Tecnológica; Universidad Nacional del Sur; Consejo Nacional de Investigaciones Científicas y Técnicas; Compute Canada","keywords":"Causal structure; Computer science; Causality (physics); Machine learning; Graph; Artificial intelligence; Time series; Causal inference; Event (particle physics); Natural language processing; Theoretical computer science; Data mining; Information retrieval; Data science; Econometrics; Mathematics","score_opus":0.04088231282822022,"score_gpt":0.3120596989749664,"score_spread":0.27117738614674614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289711940","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08524394,0.026859116,0.8720256,0.0017384414,0.0004158162,0.00038221848,0.0016752185,0.0029077677,0.00875186],"genre_scores_gemma":[0.47945553,0.020730343,0.48929155,0.00026070932,0.00086401403,0.00027192387,0.0062844846,0.00034472326,0.002496699],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99733317,0.0012053472,0.00022251904,0.0005523753,0.00059334055,0.00009319289],"domain_scores_gemma":[0.9755019,0.020630939,0.0008714209,0.0013672855,0.0014687744,0.00015970855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056018867,0.0013595296,0.0009359202,0.009900222,0.0006100981,0.0014910259,0.0012012772,0.0012745576,0.0016060194],"category_scores_gemma":[0.01963469,0.0003554676,0.0020076681,0.0076301447,0.00054983207,0.005006026,0.0008083508,0.0015921429,0.00062724337],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038530183,0.0003932214,0.010455272,0.0012504798,0.00072641415,0.00024121739,0.00062897895,0.075069845,0.0021079113,0.012579856,0.0048691253,0.8912925],"study_design_scores_gemma":[0.00008207815,0.00030076358,0.014079242,0.00029590985,0.00061818946,0.00045997833,0.00094851217,0.91997033,0.006751945,0.03274493,0.023648493,0.00009968274],"about_ca_topic_score_codex":0.002887189,"about_ca_topic_score_gemma":0.003244856,"teacher_disagreement_score":0.009900222,"about_ca_system_score_codex":0.0005814703,"about_ca_system_score_gemma":0.0008715699,"threshold_uncertainty_score":0.029625952},"labels":[],"label_agreement":null},{"id":"W4292283597","doi":"","title":"A Graph-based Approach to Cross-language Multi-document Summarization","year":2011,"lang":"en","type":"other","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Automatic summarization; Computer science; Graph; Multi-document summarization; Natural language processing; Artificial intelligence; Information retrieval; Theoretical computer science","score_opus":0.018623722499312958,"score_gpt":0.2432440398911549,"score_spread":0.22462031739184193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4292283597","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00399702,0.0006982464,0.98841304,0.00027111624,0.00009833358,0.0001967937,0.00068268913,0.0042363023,0.0014064674],"genre_scores_gemma":[0.09529789,0.00080170465,0.89332247,0.00020913096,0.0002095488,0.00030710254,0.0041725934,0.0007277076,0.004951867],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99863535,0.000556952,0.000104834675,0.00036836878,0.00028432792,0.00005011464],"domain_scores_gemma":[0.9973744,0.0010118793,0.00028586385,0.00043924578,0.00081432087,0.00007430752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012659036,0.0015414555,0.0007955647,0.0048847655,0.00085576694,0.0013889644,0.001310906,0.0011231621,0.0028769579],"category_scores_gemma":[0.0038305891,0.0004017293,0.0014984637,0.00391973,0.0005183981,0.0019483968,0.0010111695,0.0011661741,0.0016538232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021477949,0.00021875616,0.0012551082,0.0010690168,0.00058567035,0.00034571462,0.00067643815,0.066312596,0.05228786,0.021176755,0.023427518,0.83242977],"study_design_scores_gemma":[0.000091753966,0.00040516094,0.0030490255,0.00009822617,0.0006789043,0.0005132948,0.00028700664,0.8360815,0.041356973,0.0579413,0.05934141,0.00015539139],"about_ca_topic_score_codex":0.0045260615,"about_ca_topic_score_gemma":0.009095629,"teacher_disagreement_score":0.0048847655,"about_ca_system_score_codex":0.0007360833,"about_ca_system_score_gemma":0.0008445193,"threshold_uncertainty_score":0.009624422},"labels":[],"label_agreement":null},{"id":"W4292474994","doi":"10.1007/s10489-022-04052-8","title":"Transformer models used for text-based question answering systems","year":2022,"lang":"en","type":"article","venue":"Applied Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":149,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Computer science; Question answering; Transformer; Architecture; Encoder; Artificial intelligence; Natural language; Language model; Natural language processing; Information retrieval","score_opus":0.041733323832580896,"score_gpt":0.26835505980142915,"score_spread":0.22662173596884827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4292474994","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016526427,0.00062830606,0.97343713,0.0004702437,0.0001234436,0.00018682318,0.0016020107,0.004696611,0.00232909],"genre_scores_gemma":[0.6895502,0.0011944972,0.29483828,0.00032127235,0.0001816712,0.00043925375,0.0061330497,0.0010247068,0.00631717],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982698,0.0006492281,0.00017909879,0.0003843572,0.00037462736,0.00014284],"domain_scores_gemma":[0.9957489,0.0028480224,0.00015076867,0.0004486002,0.0007008784,0.00010291369],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002516646,0.00075410056,0.0009683366,0.0021582074,0.0007081936,0.0024030188,0.0016395968,0.0013236204,0.0070799105],"category_scores_gemma":[0.011746726,0.0006116209,0.0017837461,0.0024304446,0.00050614454,0.0049402327,0.0012845352,0.0017838546,0.0035243393],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015270086,0.00032118286,0.004468128,0.0009283228,0.00045759737,0.0005525508,0.0012367265,0.20983444,0.028703384,0.14949404,0.022883115,0.57959354],"study_design_scores_gemma":[0.000035106714,0.000044454828,0.0003502911,0.000022484395,0.0001077102,0.00013343364,0.00006324702,0.93359107,0.008546861,0.05059014,0.006490601,0.000024598448],"about_ca_topic_score_codex":0.008850994,"about_ca_topic_score_gemma":0.008990691,"teacher_disagreement_score":0.008850994,"about_ca_system_score_codex":0.0014777221,"about_ca_system_score_gemma":0.0014723354,"threshold_uncertainty_score":0.02368462},"labels":[],"label_agreement":null},{"id":"W4292794274","doi":"10.23919/annsim55834.2022.9859486","title":"Incremental Text Clustering Algorithm For Cloud-Based Data Management In Scientific Research Papers","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Ryerson University","keywords":"Cluster analysis; Computer science; Word2vec; Data mining; Correlation clustering; CURE data clustering algorithm; Document clustering; Canopy clustering algorithm; Conceptual clustering; Artificial intelligence; Information retrieval; Embedding","score_opus":0.14709326737545678,"score_gpt":0.3605991523184348,"score_spread":0.21350588494297804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4292794274","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055870857,0.002767268,0.9136922,0.0009255787,0.000793737,0.0011727624,0.0054643475,0.01377307,0.005540221],"genre_scores_gemma":[0.11738532,0.0008210258,0.86789805,0.00014858047,0.0002923776,0.00059998396,0.007746194,0.0003938611,0.004714579],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997792,0.00018583216,0.00026500758,0.00073002867,0.0008164771,0.00021062444],"domain_scores_gemma":[0.99616027,0.0005679569,0.0004961581,0.00063941634,0.0018196048,0.00031657508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016601755,0.0010842321,0.001281048,0.009806645,0.0019083258,0.0030079973,0.0024180976,0.0009769337,0.0035383957],"category_scores_gemma":[0.0059100436,0.00046236772,0.0015619663,0.0108222915,0.00044905598,0.0027859677,0.0019031932,0.00089930376,0.0039650677],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047649938,0.00038673787,0.007926409,0.00047779936,0.0001929457,0.00032320584,0.00055327977,0.021165367,0.019288601,0.0075107636,0.029898548,0.9117998],"study_design_scores_gemma":[0.00027374557,0.00036047387,0.011387377,0.00013063812,0.0002783183,0.0010681257,0.0013297211,0.8391948,0.044261333,0.02916033,0.07239776,0.00015741984],"about_ca_topic_score_codex":0.006904656,"about_ca_topic_score_gemma":0.009203685,"teacher_disagreement_score":0.009806645,"about_ca_system_score_codex":0.0016661619,"about_ca_system_score_gemma":0.004139427,"threshold_uncertainty_score":0.013728917},"labels":[],"label_agreement":null},{"id":"W4293024028","doi":"10.1145/3487553.3524703","title":"Biomedical Word Sense Disambiguation with Contextualized Representation Learning","year":2022,"lang":"en","type":"article","venue":"Companion Proceedings of the Web Conference 2022","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; SemEval; Word (group theory); Representation (politics); Knowledge base; Context (archaeology); Novelty; Meaning (existential); Task (project management); Linguistics","score_opus":0.03244452527112487,"score_gpt":0.2590895206372182,"score_spread":0.2266449953660933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293024028","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009732408,0.0020452093,0.98231494,0.0004520378,0.000238645,0.00013134311,0.000739966,0.0030694455,0.0012760051],"genre_scores_gemma":[0.21172553,0.0019657456,0.7787357,0.0006193603,0.00047612074,0.0004193488,0.0042803898,0.00028867662,0.0014890875],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977488,0.0007186003,0.00024795518,0.0008347162,0.0003268657,0.0001230092],"domain_scores_gemma":[0.9981002,0.0008819738,0.000220977,0.00043267405,0.00029903208,0.00006521499],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017439756,0.0014126672,0.0012694921,0.005011525,0.00075061724,0.0016891854,0.0014036761,0.0014720822,0.0020366325],"category_scores_gemma":[0.00686886,0.00043596147,0.0019878424,0.004748489,0.0008229469,0.0036912286,0.002949344,0.0015651709,0.0014841916],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002653572,0.00021080134,0.0018503051,0.0006221392,0.00033329532,0.00040410197,0.00072486047,0.046096478,0.012427297,0.027817179,0.015668496,0.8935798],"study_design_scores_gemma":[0.00011736105,0.00022930784,0.0019414057,0.0002589011,0.0003434265,0.00074128504,0.0006469955,0.7006834,0.017007062,0.2351563,0.042723183,0.00015134542],"about_ca_topic_score_codex":0.001714627,"about_ca_topic_score_gemma":0.0026786153,"teacher_disagreement_score":0.005011525,"about_ca_system_score_codex":0.0005989405,"about_ca_system_score_gemma":0.0015603862,"threshold_uncertainty_score":0.0092231035},"labels":[],"label_agreement":null},{"id":"W4293083978","doi":"10.1145/3533020","title":"Learning Implicit and Explicit Multi-task Interactions for Information Extraction","year":2022,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"National Key Research and Development Program of China","keywords":"Computer science; Multi-task learning; Generalization; Leverage (statistics); Artificial intelligence; Task (project management); Machine learning; Sequence learning; Unobservable; Representation (politics)","score_opus":0.027401701728496257,"score_gpt":0.27734139878757036,"score_spread":0.2499396970590741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293083978","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0320079,0.0006685317,0.96343565,0.00039342948,0.00006142252,0.00008167795,0.00014762077,0.0012254299,0.001978412],"genre_scores_gemma":[0.72526646,0.000668671,0.26714063,0.00043673426,0.00012509961,0.00029508534,0.0010482906,0.0002221979,0.0047968756],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988242,0.0003354883,0.00006866879,0.00042599565,0.00020839904,0.00013730608],"domain_scores_gemma":[0.9974469,0.0014162936,0.00024541924,0.0004675387,0.0002993586,0.0001244251],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020461048,0.0017822252,0.0009503917,0.00058855885,0.0006145194,0.0009032622,0.0020282045,0.0013061129,0.0016333252],"category_scores_gemma":[0.007025431,0.000654783,0.001265249,0.00072418194,0.0007797832,0.0038392087,0.0023731636,0.0032743318,0.00081359525],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048410197,0.00061035994,0.005945704,0.0003428423,0.00034435978,0.00058627845,0.0006781832,0.4529666,0.03813067,0.028020948,0.006107283,0.46578267],"study_design_scores_gemma":[0.000007064311,0.000055409637,0.00048412662,0.000007127169,0.000027171589,0.0000413275,0.000012897872,0.9842481,0.004296119,0.009989522,0.00081834313,0.000012757726],"about_ca_topic_score_codex":0.0031140563,"about_ca_topic_score_gemma":0.0062022447,"teacher_disagreement_score":0.0031140563,"about_ca_system_score_codex":0.00092146604,"about_ca_system_score_gemma":0.0010324302,"threshold_uncertainty_score":0.010820985},"labels":[],"label_agreement":null},{"id":"W4293428470","doi":"10.1109/tai.2022.3201807","title":"Optimizing Multidocument Summarization by Blending Reinforcement Learning Policies","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Science Foundation","keywords":"Automatic summarization; Computer science; Reinforcement learning; Relevance (law); Redundancy (engineering); Multi-document summarization; Sentence; Artificial intelligence; Quality (philosophy); Machine learning; Information retrieval; Natural language processing","score_opus":0.04336013915883275,"score_gpt":0.2878445156778042,"score_spread":0.24448437651897148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293428470","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03227767,0.0006708348,0.96387655,0.00027726684,0.000047911595,0.0001356326,0.00006570057,0.0016906196,0.00095777237],"genre_scores_gemma":[0.5039702,0.00055490475,0.4901277,0.00033246772,0.00015041335,0.00043962654,0.00048589066,0.00034899803,0.0035896965],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998868,0.00039373664,0.00009743156,0.00033017236,0.0002224067,0.00008828404],"domain_scores_gemma":[0.9965205,0.0023017188,0.0003909412,0.0001988337,0.00043422362,0.0001538462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027194498,0.0013986988,0.0016371656,0.0010393724,0.00045985638,0.0011173945,0.001197242,0.0016984983,0.0015315848],"category_scores_gemma":[0.0073111304,0.00059979706,0.0007485653,0.000982467,0.00065049686,0.002248365,0.0011130071,0.001570469,0.00058190664],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021691134,0.00029189585,0.0009648538,0.00026145144,0.00010941952,0.00012892154,0.00019338084,0.7688207,0.013352959,0.0035206943,0.0015333316,0.21060543],"study_design_scores_gemma":[0.000026342523,0.00007798693,0.000103951286,0.000007370599,0.000017192486,0.000014675615,0.000014649397,0.99377936,0.003180128,0.0022056536,0.0005644914,0.000008260806],"about_ca_topic_score_codex":0.0025921885,"about_ca_topic_score_gemma":0.0032374754,"teacher_disagreement_score":0.0027194498,"about_ca_system_score_codex":0.0012866404,"about_ca_system_score_gemma":0.0013634157,"threshold_uncertainty_score":0.014382005},"labels":[],"label_agreement":null},{"id":"W4293918660","doi":"10.1101/2022.08.30.22279318","title":"Evaluating Progress in Automatic Chest X-Ray Radiology Report Generation","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institute of Biomedical Imaging and Bioengineering; National Institutes of Health","keywords":"Workflow; Computer science; Metric (unit); Rank (graph theory); Quality (philosophy); Artificial intelligence; Medical imaging; Identification (biology); Contrast (vision); Medical physics; Information retrieval; Machine learning; Natural language processing; Data science; Radiology; Medicine; Database","score_opus":0.11079709504860116,"score_gpt":0.3653045299124789,"score_spread":0.2545074348638777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293918660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7436329,0.00883617,0.19346802,0.0023716986,0.0007576866,0.0010357909,0.0043566027,0.039405443,0.0061356705],"genre_scores_gemma":[0.823254,0.0006223581,0.16358456,0.00020982424,0.0001813797,0.00022514311,0.010141209,0.00071866,0.0010628433],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.942482,0.03560451,0.0048637013,0.0047534127,0.011338579,0.00095773063],"domain_scores_gemma":[0.74609125,0.17281435,0.016832113,0.022498751,0.037773978,0.0039896383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.055031557,0.002305701,0.0014587917,0.008304799,0.00080574997,0.0047518904,0.003783455,0.002289626,0.0014351323],"category_scores_gemma":[0.17196336,0.00063344353,0.0010482915,0.003498847,0.0010622567,0.0035064663,0.002888364,0.0016656725,0.0012173356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029892414,0.0018493577,0.08658924,0.0021270846,0.0009067888,0.0003046712,0.0017397006,0.17534393,0.016336525,0.0040256414,0.02612022,0.6816676],"study_design_scores_gemma":[0.00043987538,0.0026295558,0.033110287,0.00021045677,0.00028716034,0.0004472501,0.0007691951,0.8904234,0.056342375,0.003770377,0.011338663,0.00023130691],"about_ca_topic_score_codex":0.005463732,"about_ca_topic_score_gemma":0.0035406828,"teacher_disagreement_score":0.055031557,"about_ca_system_score_codex":0.0022708317,"about_ca_system_score_gemma":0.002380548,"threshold_uncertainty_score":0.29103816},"labels":[],"label_agreement":null},{"id":"W4293920114","doi":"10.54097/hset.v12i.1368","title":"The Advance of Deep Learning Based Named Entity Recognition","year":2022,"lang":"en","type":"article","venue":"Highlights in Science Engineering and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Named-entity recognition; Computer science; Artificial intelligence; Deep learning; Entity linking; Named entity; Variety (cybernetics); Natural language processing; Artificial neural network; Machine learning; Knowledge base; Task (project management)","score_opus":0.006035953314598795,"score_gpt":0.19968078156759286,"score_spread":0.19364482825299406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293920114","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009178638,0.014687633,0.9622498,0.0032521815,0.00056999305,0.000036776135,0.00076627,0.0025554944,0.0067031863],"genre_scores_gemma":[0.43750697,0.038143355,0.49681893,0.0021917415,0.0014952015,0.00014379165,0.006930629,0.00047149355,0.016297942],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988726,0.0003223634,0.000086227,0.0003263279,0.0003084664,0.00008395761],"domain_scores_gemma":[0.9984376,0.0005935486,0.00012334666,0.00036233172,0.00042961448,0.000053583455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018196009,0.00086578337,0.0010452734,0.0013746696,0.000313959,0.0017308736,0.001684548,0.0011252573,0.0023594315],"category_scores_gemma":[0.003685005,0.00037016897,0.0008984653,0.0024016509,0.00071683637,0.007279919,0.001242092,0.002474042,0.001874006],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010893508,0.00010054949,0.002755201,0.0006366307,0.0001902525,0.00016815362,0.00020005753,0.060848698,0.0103332745,0.08964533,0.028642783,0.8063701],"study_design_scores_gemma":[0.000008851831,0.00004396535,0.0011872676,0.000102542166,0.00007921608,0.00015472702,0.00007264009,0.8426191,0.011677083,0.08106147,0.06293355,0.000059607886],"about_ca_topic_score_codex":0.0037902868,"about_ca_topic_score_gemma":0.002998611,"teacher_disagreement_score":0.0037902868,"about_ca_system_score_codex":0.0011010362,"about_ca_system_score_gemma":0.00094721554,"threshold_uncertainty_score":0.009623051},"labels":[],"label_agreement":null},{"id":"W4295005888","doi":"","title":"Semantics Altering Modifications for Evaluating Comprehension in Machine Reading","year":2020,"lang":"en","type":"preprint","venue":"Research Explorer (The University of Manchester)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Computer science; Semantics (computer science); Natural language processing; Artificial intelligence; Process (computing); Sentence; Comprehension; Domain (mathematical analysis); Reading (process); Machine learning; Programming language; Linguistics","score_opus":0.3797729647456968,"score_gpt":0.39520795626857463,"score_spread":0.015434991522877861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295005888","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6065909,0.0049354676,0.3548376,0.0013564177,0.00031908485,0.00079513615,0.002819995,0.014764455,0.013580915],"genre_scores_gemma":[0.9081848,0.0004794625,0.0857188,0.00016337255,0.00006630925,0.00026303163,0.0030815734,0.00033714395,0.0017056434],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99675566,0.0016546049,0.0002707179,0.0006813459,0.00051058386,0.00012711667],"domain_scores_gemma":[0.98206466,0.013019091,0.0011573307,0.0021280947,0.0012720712,0.0003586868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059086597,0.0019948655,0.0008273886,0.0021128121,0.0004546595,0.0023197753,0.0017888049,0.0027781136,0.0037059514],"category_scores_gemma":[0.033191763,0.00048121216,0.0008153023,0.001361695,0.0011297641,0.005175901,0.0018483493,0.0030832158,0.0014116088],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001574417,0.0008211143,0.032986388,0.0022362706,0.0008491631,0.00038360304,0.001777432,0.2371497,0.06585747,0.004985202,0.009691427,0.6416878],"study_design_scores_gemma":[0.000092596936,0.0012357745,0.017528927,0.00012720404,0.00020485229,0.0002588336,0.0005481676,0.90336007,0.0591217,0.013094951,0.0043217815,0.000105075385],"about_ca_topic_score_codex":0.0024382712,"about_ca_topic_score_gemma":0.0039717644,"teacher_disagreement_score":0.0059086597,"about_ca_system_score_codex":0.0011536666,"about_ca_system_score_gemma":0.00086369016,"threshold_uncertainty_score":0.031248331},"labels":[],"label_agreement":null},{"id":"W4295308551","doi":"10.1109/tai.2022.3205567","title":"DReD–A Descriptive Relation Dataset for Expanding Relation Extraction","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Relationship extraction; Computer science; Relation (database); Benchmark (surveying); Sentence; Natural language processing; Task (project management); Artificial intelligence; Code (set theory); Set (abstract data type); Information retrieval; Data mining","score_opus":0.12257939294696753,"score_gpt":0.3446536909603573,"score_spread":0.2220742980133898,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295308551","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033243768,0.0016866183,0.026762504,0.0010968134,0.00028823185,0.00072384265,0.90022415,0.024045974,0.011928151],"genre_scores_gemma":[0.014108355,0.00018128491,0.03365224,0.0001979682,0.000034302153,0.00044931477,0.94937027,0.00034835326,0.0016579579],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9948571,0.0008882221,0.00089281827,0.0013973794,0.0017106388,0.0002539225],"domain_scores_gemma":[0.9899603,0.003286449,0.0009416822,0.0032909026,0.0019818528,0.00053870934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026126252,0.0021714277,0.0010549497,0.0076232087,0.0020793353,0.0017058118,0.0039262515,0.0028929203,0.007899724],"category_scores_gemma":[0.01163561,0.0006885159,0.0017722541,0.0069567035,0.00082777586,0.004588405,0.0024953256,0.0026596123,0.010741408],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000539166,0.00076579105,0.012468941,0.003047232,0.00019773708,0.001017595,0.00072736334,0.005772845,0.0135699455,0.011673534,0.85334176,0.09687813],"study_design_scores_gemma":[0.0003868384,0.0003037743,0.028846888,0.0004218681,0.00013666676,0.0017712469,0.00096688996,0.04777291,0.024488121,0.010538053,0.8841552,0.00021151704],"about_ca_topic_score_codex":0.015619209,"about_ca_topic_score_gemma":0.031090748,"teacher_disagreement_score":0.015619209,"about_ca_system_score_codex":0.0022306787,"about_ca_system_score_gemma":0.0029496413,"threshold_uncertainty_score":0.031056583},"labels":[],"label_agreement":null},{"id":"W4295940462","doi":"10.1007/978-3-031-16270-1_17","title":"A Novel Hybrid Framework to Enhance Zero-shot Classification","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wilfrid Laurier University","funders":"","keywords":"Computer science; Exploit; Transformer; Artificial intelligence; Categorical variable; Language model; Data mining; Machine learning; Natural language processing","score_opus":0.03860569090495036,"score_gpt":0.28756313973566006,"score_spread":0.2489574488307097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295940462","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010461488,0.0014718267,0.9800046,0.00017034805,0.00026795318,0.000100937694,0.0003587277,0.004506877,0.0026572957],"genre_scores_gemma":[0.18827544,0.001210637,0.78000575,0.0005269783,0.0005985015,0.00026909256,0.0040035355,0.0010934114,0.024016634],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985493,0.0002380828,0.00005851817,0.00041390955,0.0005391988,0.00020091436],"domain_scores_gemma":[0.99904495,0.00023608778,0.000035263954,0.00020317538,0.00038026783,0.000100212594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012164923,0.001371127,0.0020894003,0.002619838,0.0010025721,0.0021472413,0.0029006647,0.0020681059,0.006065224],"category_scores_gemma":[0.0018378516,0.00045447142,0.0013366676,0.0023296946,0.00053132983,0.0029625667,0.0027417499,0.0018991857,0.0046547893],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003518572,0.00043963385,0.00068167294,0.00019619713,0.00013982905,0.00009320763,0.000115579154,0.012192694,0.043263003,0.0069799805,0.016389288,0.91915697],"study_design_scores_gemma":[0.00002858176,0.00014306686,0.0009494537,0.000028918066,0.00010594649,0.00022565627,0.00006033023,0.9538683,0.0213451,0.011162414,0.012036475,0.000045765264],"about_ca_topic_score_codex":0.008856818,"about_ca_topic_score_gemma":0.01899955,"teacher_disagreement_score":0.008856818,"about_ca_system_score_codex":0.00067359547,"about_ca_system_score_gemma":0.0014255132,"threshold_uncertainty_score":0.020290196},"labels":[],"label_agreement":null},{"id":"W4296711106","doi":"10.1162/tacl_a_00506","title":"Evaluating Attribution in Dialogue Systems: The BEGIN Benchmark","year":2022,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Spurious relationship; Benchmark (surveying); Attribution; Artificial intelligence; Natural language processing; Grounded theory; Data science; Machine learning; Qualitative research","score_opus":0.04673618613088275,"score_gpt":0.30956918011848517,"score_spread":0.2628329939876024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296711106","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73103005,0.01103184,0.15626968,0.0015897434,0.0015991917,0.0019223475,0.029711638,0.042405706,0.024439735],"genre_scores_gemma":[0.8645595,0.0006093174,0.06819958,0.000445379,0.0002618006,0.0011525924,0.057925884,0.0015903887,0.005255618],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9777314,0.014776288,0.0011039358,0.0031783422,0.0024833926,0.0007266432],"domain_scores_gemma":[0.96450174,0.021673985,0.0016378951,0.0056798817,0.004821747,0.001684825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01285615,0.0028834313,0.0013141662,0.003055743,0.0012040054,0.0028749462,0.0025310463,0.0026465117,0.0038888773],"category_scores_gemma":[0.05178178,0.0005784282,0.0010920241,0.0018417037,0.0014094716,0.003660867,0.0047868458,0.0030145852,0.0033331919],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007725429,0.004401443,0.05878584,0.0056298,0.0015873449,0.000884786,0.005734972,0.21265575,0.030051304,0.008207325,0.12640616,0.5379299],"study_design_scores_gemma":[0.0010292466,0.0050476273,0.054935865,0.0007914133,0.00039167807,0.00075554755,0.0041089244,0.7847366,0.059958328,0.022159478,0.065623224,0.00046207898],"about_ca_topic_score_codex":0.005043851,"about_ca_topic_score_gemma":0.0072882837,"teacher_disagreement_score":0.01285615,"about_ca_system_score_codex":0.001877344,"about_ca_system_score_gemma":0.001374885,"threshold_uncertainty_score":0.06799066},"labels":[],"label_agreement":null},{"id":"W4296878065","doi":"10.1007/978-3-031-17120-8_31","title":"Coarse-to-Fine Retriever for Better Open-Domain Question Answering","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Labrador Retriever; Computer science; Domain (mathematical analysis); Representation (politics); Construct (python library); Information retrieval; Artificial intelligence; Reading (process); Natural language processing; Linguistics; Mathematics; Programming language","score_opus":0.02721815459113205,"score_gpt":0.2740297315924848,"score_spread":0.24681157700135273,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296878065","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029288696,0.003209458,0.7698499,0.0014603212,0.0010904229,0.0007435098,0.008192949,0.17266354,0.01350118],"genre_scores_gemma":[0.13740303,0.0010532188,0.7980401,0.0014992324,0.0007878595,0.00042299,0.03068357,0.008193841,0.02191618],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974093,0.0006308859,0.00032513094,0.0008009514,0.0005555954,0.00027821117],"domain_scores_gemma":[0.9935096,0.0022347271,0.00014121275,0.0027345896,0.0010487272,0.00033112368],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023971198,0.0026155345,0.003106172,0.0041977228,0.0012356173,0.0037160763,0.0031142368,0.0031561207,0.05008301],"category_scores_gemma":[0.0103554195,0.0010560554,0.0021968747,0.0032629922,0.0008846363,0.010059564,0.0064116353,0.0036504588,0.044418924],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008201114,0.0006521488,0.0011431768,0.00091802806,0.00013270564,0.00036380443,0.0004331142,0.003504741,0.05533867,0.010071245,0.15465201,0.7719702],"study_design_scores_gemma":[0.0005226311,0.0008632884,0.0034450558,0.00027336096,0.00042345797,0.0020499628,0.0012773369,0.5554572,0.12901054,0.103395514,0.202963,0.00031859684],"about_ca_topic_score_codex":0.00366304,"about_ca_topic_score_gemma":0.0072179567,"teacher_disagreement_score":0.05008301,"about_ca_system_score_codex":0.0008302132,"about_ca_system_score_gemma":0.0016792286,"threshold_uncertainty_score":0.16754436},"labels":[],"label_agreement":null},{"id":"W4296941844","doi":"10.1101/2022.09.22.22280246","title":"Large-Scale Application of Named Entity Recognition to Biomedicine and Epidemiology","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Public Health Ontario; University of Toronto","funders":"Institute of Health Services and Policy Research; Canadian Institutes of Health Research","keywords":"Computer science; Named-entity recognition; Inference; Preprocessor; Biomedicine; Parsing; Data science; Artificial intelligence; Declaration; Machine learning; Natural language processing; Bioinformatics","score_opus":0.05218879559191653,"score_gpt":0.3170996376059271,"score_spread":0.26491084201401055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296941844","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07037706,0.009428048,0.6308399,0.008851811,0.0016091961,0.0016502204,0.16487454,0.10316075,0.009208492],"genre_scores_gemma":[0.25926602,0.0029669318,0.46861318,0.002349809,0.00062525645,0.0008762355,0.26026174,0.00091366755,0.004127138],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99552554,0.0015400342,0.00065072416,0.0015515652,0.0005688088,0.0001633349],"domain_scores_gemma":[0.98737794,0.0075454623,0.0008590726,0.0021454603,0.0015961934,0.00047580147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006213411,0.0016091374,0.0010478691,0.0058924314,0.0009222053,0.0016092269,0.0025762047,0.0017675355,0.004345717],"category_scores_gemma":[0.014764266,0.0005398495,0.0022666932,0.004145098,0.0006217732,0.0030821855,0.0028608444,0.0017279718,0.0038296082],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075302046,0.0007170572,0.050141852,0.0046735713,0.0010904502,0.003029855,0.0006743283,0.074132584,0.017968554,0.0076782787,0.20147474,0.63766575],"study_design_scores_gemma":[0.00016898368,0.00027866173,0.053058416,0.00068187795,0.00068582915,0.0032248178,0.0008147771,0.6561749,0.050535332,0.03851391,0.19556862,0.00029390273],"about_ca_topic_score_codex":0.009229947,"about_ca_topic_score_gemma":0.010723855,"teacher_disagreement_score":0.009229947,"about_ca_system_score_codex":0.001360652,"about_ca_system_score_gemma":0.0025222672,"threshold_uncertainty_score":0.03286004},"labels":[],"label_agreement":null},{"id":"W4297796135","doi":"10.1145/3548785.3548789","title":"An Online MCQ sub-system for CrsMgr","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Transfer of learning; Focus (optics); Human–computer interaction; Online learning; Similarity (geometry); Artificial intelligence; Multimedia","score_opus":0.057336339571853454,"score_gpt":0.2722390468146272,"score_spread":0.21490270724277372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297796135","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012069047,0.0002775972,0.23588306,0.0005108307,0.0002018009,0.0016937802,0.017770508,0.69511455,0.036478844],"genre_scores_gemma":[0.32004964,0.00034488493,0.40694487,0.002388006,0.00055648095,0.0031017626,0.09138237,0.07429503,0.10093697],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983193,0.00024237798,0.00014046585,0.00061448466,0.0005414126,0.00014203774],"domain_scores_gemma":[0.996416,0.0009274092,0.00017292709,0.0012679971,0.0009309866,0.00028462659],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022487843,0.0010560931,0.00097401225,0.0024966756,0.00068353105,0.0015374874,0.0026335383,0.0012418478,0.13349468],"category_scores_gemma":[0.00718517,0.00051029125,0.00065588887,0.0013209129,0.00037855125,0.00236241,0.0021412403,0.0008285077,0.061133154],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015230069,0.00034654717,0.004141459,0.00070461107,0.000117973774,0.0007035392,0.000824023,0.0026755657,0.03896068,0.0071798814,0.39756885,0.5452538],"study_design_scores_gemma":[0.0007405513,0.00048450078,0.008947487,0.00018385585,0.00010584539,0.001038439,0.00030600606,0.13054232,0.057141922,0.008060353,0.79219604,0.0002528107],"about_ca_topic_score_codex":0.009346045,"about_ca_topic_score_gemma":0.0045297733,"teacher_disagreement_score":0.13349468,"about_ca_system_score_codex":0.0013644231,"about_ca_system_score_gemma":0.0011200919,"threshold_uncertainty_score":0.44658417},"labels":[],"label_agreement":null},{"id":"W4297941332","doi":"10.1109/compsac54236.2022.00210","title":"A Question-Answering System on COVID-19 Scientific Literature","year":2022,"lang":"en","type":"article","venue":"2022 IEEE 46th Annual Computers, Software, and Applications Conference (COMPSAC)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Public Health; Public Health Ontario; University of Toronto","funders":"","keywords":"Coronavirus disease 2019 (COVID-19); Computer science; Pipeline (software); Gold standard (test); Question answering; Information retrieval; Data science; Natural language processing; Infectious disease (medical specialty); Statistics","score_opus":0.022063189100869015,"score_gpt":0.26210149271896377,"score_spread":0.24003830361809475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297941332","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07692255,0.008625188,0.21951668,0.008412114,0.0020399536,0.004567158,0.3562673,0.29027554,0.033373464],"genre_scores_gemma":[0.09567554,0.0011904757,0.37745667,0.0028689504,0.00047559,0.0018783312,0.5101914,0.0016778922,0.008585172],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99683654,0.0007548338,0.00039504262,0.0012447349,0.0006055371,0.00016332912],"domain_scores_gemma":[0.9923625,0.0039660493,0.0004232789,0.0008076705,0.0019892952,0.00045113743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004641403,0.0025014551,0.0014137183,0.012594836,0.0014689909,0.0024785732,0.002831983,0.0032151574,0.021496478],"category_scores_gemma":[0.018097544,0.0006287172,0.0016004583,0.0045062895,0.00058919867,0.0069340724,0.0054768766,0.0019259097,0.014729132],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008178961,0.0007577425,0.007935599,0.005743162,0.00024161482,0.0012460287,0.0016843082,0.0056125214,0.021723848,0.009232783,0.56757766,0.37742677],"study_design_scores_gemma":[0.0005874299,0.00082663604,0.012166025,0.0011674564,0.00037593322,0.0015180034,0.0027396653,0.3328622,0.038236555,0.036247157,0.5729831,0.00028992438],"about_ca_topic_score_codex":0.008714993,"about_ca_topic_score_gemma":0.01385814,"teacher_disagreement_score":0.021496478,"about_ca_system_score_codex":0.0025560132,"about_ca_system_score_gemma":0.0033225084,"threshold_uncertainty_score":0.071912885},"labels":[],"label_agreement":null},{"id":"W4297990402","doi":"10.18280/ria.360409","title":"Deep Named Entity Recognition in Hindi Using Neural Networks","year":2022,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Named-entity recognition; Computer science; Natural language processing; Artificial intelligence; Hindi; Phrase; Task (project management); Deep learning; Word (group theory); Autoencoder; Architecture; Named entity; Recurrent neural network; Entity linking; Artificial neural network; Linguistics; Knowledge base","score_opus":0.07522228084441177,"score_gpt":0.27589825618913943,"score_spread":0.20067597534472767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297990402","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54742026,0.0026440881,0.4175009,0.00089046056,0.00045098166,0.00020303165,0.0037737587,0.011942391,0.015174144],"genre_scores_gemma":[0.8556451,0.00045733663,0.13224751,0.00016502017,0.000058838974,0.00008776233,0.004819245,0.00007333495,0.006445702],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99978894,0.00005349912,0.000012887142,0.000073198855,0.00003764984,0.000033837176],"domain_scores_gemma":[0.99958533,0.00018337835,0.000029393717,0.00007013469,0.00010958985,0.000022148348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005001039,0.00038412743,0.0003030456,0.00038563984,0.00029873953,0.00053005613,0.0005947633,0.00042185694,0.0016503187],"category_scores_gemma":[0.0011147191,0.00015894175,0.00023853376,0.00059457513,0.00021996365,0.0014387093,0.00056142226,0.00066879985,0.00104977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005423036,0.000290117,0.004192968,0.00025678638,0.00008694716,0.00066965556,0.00056000607,0.17333007,0.05019774,0.009965386,0.02017907,0.739729],"study_design_scores_gemma":[0.000014685956,0.00008847201,0.0033291073,0.000015097759,0.000014073598,0.00009617981,0.00009341459,0.9688945,0.015690027,0.0059346566,0.005807071,0.000022745293],"about_ca_topic_score_codex":0.010196698,"about_ca_topic_score_gemma":0.013973919,"teacher_disagreement_score":0.010196698,"about_ca_system_score_codex":0.0006621343,"about_ca_system_score_gemma":0.0003906696,"threshold_uncertainty_score":0.020274699},"labels":[],"label_agreement":null},{"id":"W4298095693","doi":"10.48550/arxiv.2109.10739","title":"Predicting Efficiency/Effectiveness Trade-offs for Dense vs. Sparse\\n Retrieval Strategy Selection","year":2021,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Leverage (statistics); Embedding; Classifier (UML); Artificial intelligence","score_opus":0.08046011134342035,"score_gpt":0.21064954058147806,"score_spread":0.1301894292380577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4298095693","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84877425,0.015554977,0.11662633,0.0013608325,0.0001363561,0.00053806644,0.0023154102,0.004643183,0.010050579],"genre_scores_gemma":[0.9278104,0.0023862384,0.06292322,0.00029357377,0.00012956312,0.00016784185,0.0031691429,0.00027574893,0.0028442903],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99802977,0.00043031143,0.0002633297,0.00046857988,0.0005337568,0.00027436853],"domain_scores_gemma":[0.9912548,0.006733705,0.00043763072,0.00075289165,0.0006131622,0.00020774706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028597154,0.0010128412,0.0011963175,0.002094495,0.0003526053,0.0016132524,0.0009994479,0.0013396267,0.0019436014],"category_scores_gemma":[0.015428449,0.00036560168,0.0006359233,0.0016678498,0.0006057246,0.0031317691,0.00061734684,0.0008028701,0.00193682],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029618742,0.0014080022,0.048233226,0.0017712762,0.00047942475,0.000422767,0.0002667518,0.18080369,0.041740246,0.0039303196,0.021839155,0.6961432],"study_design_scores_gemma":[0.00015926918,0.00094273506,0.007227033,0.000048787027,0.00020511058,0.00048241485,0.00019364202,0.96575314,0.018159209,0.0038148821,0.0029665793,0.000047244495],"about_ca_topic_score_codex":0.00551549,"about_ca_topic_score_gemma":0.008560661,"teacher_disagreement_score":0.00551549,"about_ca_system_score_codex":0.0008640054,"about_ca_system_score_gemma":0.0010426085,"threshold_uncertainty_score":0.015123785},"labels":[],"label_agreement":null},{"id":"W4299441762","doi":"","title":"Predicting the Semantic Textual Similarity with Siamese CNN and LSTM","year":2018,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Semantic similarity; Artificial intelligence; Natural language processing; Similarity (geometry); Semantics (computer science); Information retrieval","score_opus":0.012785283939800038,"score_gpt":0.21597728854243023,"score_spread":0.2031920046026302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4299441762","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77526975,0.0051615043,0.19049028,0.002114902,0.0016883545,0.00024718343,0.00762077,0.0068786065,0.010528714],"genre_scores_gemma":[0.95006466,0.0006759049,0.03401659,0.00025061608,0.0005873215,0.00007076066,0.0068094693,0.00015522493,0.0073695257],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995751,0.00008040378,0.000031741576,0.00017793672,0.00006944201,0.00006537522],"domain_scores_gemma":[0.9990381,0.0004676713,0.00007737797,0.00008694115,0.00024799776,0.00008189493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006522637,0.0011021008,0.0006857151,0.0017993341,0.0003639058,0.0010099246,0.000705172,0.0014307778,0.003825879],"category_scores_gemma":[0.002871523,0.00022174513,0.0009004746,0.0014834089,0.0002873068,0.0019698387,0.0006642412,0.0010384446,0.0024415036],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025109244,0.0009830606,0.017534178,0.0007534835,0.00046887613,0.0007500483,0.00018372321,0.05568447,0.07798618,0.0034176197,0.036399983,0.8033274],"study_design_scores_gemma":[0.00003778651,0.00018514862,0.0053797397,0.00001982187,0.00010218737,0.00013892086,0.00007067493,0.9795639,0.009819226,0.003165942,0.00149999,0.000016747654],"about_ca_topic_score_codex":0.0066747996,"about_ca_topic_score_gemma":0.011661109,"teacher_disagreement_score":0.0066747996,"about_ca_system_score_codex":0.00074648595,"about_ca_system_score_gemma":0.0006534203,"threshold_uncertainty_score":0.013271868},"labels":[],"label_agreement":null},{"id":"W4300427152","doi":"10.1117/12.2641031","title":"Self-attention on RNN-based text classification","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Recurrent neural network; Artificial intelligence; Artificial neural network","score_opus":0.03357897251242559,"score_gpt":0.24794164988067163,"score_spread":0.21436267736824605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4300427152","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49951345,0.0074623297,0.46970758,0.0024164878,0.00088630914,0.00023488663,0.0009563302,0.008050178,0.010772437],"genre_scores_gemma":[0.9410233,0.0006741233,0.050836083,0.00034122364,0.00023756138,0.00006841422,0.00089343195,0.00017550038,0.005750401],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988803,0.0004223255,0.00007468136,0.00031754427,0.00016310585,0.00014203016],"domain_scores_gemma":[0.99630225,0.0022444276,0.00023801146,0.00039355087,0.0006861764,0.00013561355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00417726,0.0009961013,0.00081543124,0.0014618573,0.0003865007,0.00092021585,0.0013397052,0.0010936293,0.0018706752],"category_scores_gemma":[0.00841114,0.00031941576,0.0006812899,0.00092938007,0.00044563712,0.0028327773,0.0011207969,0.0012691492,0.0008450937],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009979036,0.00046008115,0.008700261,0.0004632673,0.00037120082,0.0002937771,0.0004122088,0.4141909,0.01960504,0.007114439,0.01017785,0.537213],"study_design_scores_gemma":[0.000005175736,0.00004637158,0.0009997041,0.000011114965,0.000015852092,0.000020928213,0.000014017374,0.9939937,0.0027201003,0.0017491316,0.00041663498,0.000007323494],"about_ca_topic_score_codex":0.011370175,"about_ca_topic_score_gemma":0.0108011635,"teacher_disagreement_score":0.011370175,"about_ca_system_score_codex":0.0014139984,"about_ca_system_score_gemma":0.00060190965,"threshold_uncertainty_score":0.022607982},"labels":[],"label_agreement":null},{"id":"W4301054811","doi":"10.48550/arxiv.1805.04558","title":"NRC-Canada at SMM4H Shared Task: Classifying Tweets Mentioning Adverse\\n Drug Reactions and Medication Intake","year":2018,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Lexicon; Task (project management); Class (philosophy); Computer science; Domain (mathematical analysis); Variety (cybernetics); Natural language processing; Artificial intelligence; Support vector machine; Word (group theory); Social media; Sentiment analysis; World Wide Web; Mathematics; Engineering","score_opus":0.06374858205544988,"score_gpt":0.18946913828064055,"score_spread":0.12572055622519068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4301054811","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08969929,0.0027624147,0.04755015,0.014595304,0.00865768,0.003485819,0.7290923,0.055106036,0.049050957],"genre_scores_gemma":[0.10424635,0.0007502708,0.07639511,0.0021625815,0.0012824449,0.002110005,0.7420117,0.0046755173,0.06636597],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941355,0.001338551,0.00027045133,0.0012611097,0.0021107763,0.0008835663],"domain_scores_gemma":[0.9821454,0.0034807273,0.00046839713,0.0032896616,0.007856799,0.002758894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007520223,0.0034446516,0.0025307732,0.0035021119,0.0058923084,0.0040283664,0.0027842699,0.0037970233,0.029019555],"category_scores_gemma":[0.022273606,0.0008772764,0.0018259778,0.0034205744,0.0014133042,0.0028676523,0.0050260667,0.0032778506,0.024897803],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006947744,0.0002919621,0.004878797,0.0004650548,0.00013399083,0.00020116403,0.00041692905,0.001876358,0.003965717,0.0008610915,0.930045,0.05616907],"study_design_scores_gemma":[0.0009167154,0.0004691287,0.04334613,0.00028587642,0.0002448383,0.00025989485,0.0023446572,0.066331394,0.019225996,0.0055841184,0.86058223,0.00040887002],"about_ca_topic_score_codex":0.43836424,"about_ca_topic_score_gemma":0.62569183,"teacher_disagreement_score":0.43836424,"about_ca_system_score_codex":0.0068982635,"about_ca_system_score_gemma":0.018523825,"threshold_uncertainty_score":0.87162536},"labels":[],"label_agreement":null},{"id":"W4302599139","doi":"10.1007/978-1-4614-6170-8_100005","title":"Information Extraction","year":2014,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Extraction (chemistry); Information extraction; Computer science; Information retrieval; Chromatography; Chemistry","score_opus":0.021307402991167224,"score_gpt":0.2286075006550297,"score_spread":0.20730009766386245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4302599139","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027452514,0.023492897,0.6985403,0.0025258271,0.0018868707,0.0006829688,0.011118506,0.01742471,0.24158259],"genre_scores_gemma":[0.034415156,0.037793394,0.47365093,0.0014839398,0.00254824,0.0006893466,0.05089544,0.0050435937,0.39347997],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99925834,0.00011518734,0.00007286631,0.00021859966,0.0002880281,0.000046865607],"domain_scores_gemma":[0.998896,0.0004111833,0.000050415467,0.00028128852,0.0003182751,0.000042796844],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008563232,0.0021663876,0.0014621094,0.0075578024,0.0011517133,0.004256593,0.0016635038,0.0009928998,0.07937279],"category_scores_gemma":[0.0031620325,0.00084902963,0.0015622994,0.0091086,0.0005170731,0.0066449395,0.0022623646,0.0015588462,0.096333675],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029291796,0.00004217065,0.00015355137,0.0005801344,0.00003572443,0.000066732355,0.000112370595,0.0006433836,0.0030792025,0.015974049,0.15845035,0.820833],"study_design_scores_gemma":[0.0000124099415,0.000028472494,0.00075998256,0.00047435606,0.00010097147,0.0006213839,0.0001501787,0.008568574,0.011586858,0.052012384,0.92563564,0.000048820188],"about_ca_topic_score_codex":0.0014110969,"about_ca_topic_score_gemma":0.0019937172,"teacher_disagreement_score":0.07937279,"about_ca_system_score_codex":0.00076718564,"about_ca_system_score_gemma":0.0014372079,"threshold_uncertainty_score":0.26552844},"labels":[],"label_agreement":null},{"id":"W4304698333","doi":"10.1007/s10462-022-10265-7","title":"Deep learning, graph-based text representation and classification: a survey, perspectives and challenges","year":2022,"lang":"en","type":"article","venue":"Artificial Intelligence Review","topic":"Topic Modeling","field":"Computer Science","cited_by":65,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Artificial intelligence; Deep learning; Feature engineering; Recurrent neural network; Feature learning; Graph; Artificial neural network; Representation (politics); Machine learning; Natural language processing; Theoretical computer science","score_opus":0.22418225187240653,"score_gpt":0.3617897368252568,"score_spread":0.13760748495285025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4304698333","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015779963,0.5401441,0.41240463,0.017660424,0.0010678275,0.00019218556,0.0030189238,0.0018948935,0.0078370515],"genre_scores_gemma":[0.13743031,0.6576266,0.18289904,0.00263705,0.004311638,0.0003094033,0.007835978,0.0002820894,0.006667778],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992576,0.00020701363,0.00006697037,0.0001718932,0.00025149612,0.000045093264],"domain_scores_gemma":[0.99694973,0.001999097,0.00018775024,0.00018825442,0.00057073997,0.000104480205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018689913,0.000970932,0.0019236678,0.0035719732,0.00030966345,0.0023605959,0.0020009438,0.0010661397,0.0018850438],"category_scores_gemma":[0.0050130878,0.0003466739,0.0007629781,0.005667462,0.0006689742,0.005182645,0.0010780606,0.001908405,0.0015338719],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000574596,0.00015336332,0.0015871621,0.0021369949,0.000085866486,0.000017087561,0.00008679138,0.008438777,0.0010899691,0.014410738,0.027332084,0.9446037],"study_design_scores_gemma":[0.000052044925,0.00041201073,0.0070331884,0.002618586,0.00036085222,0.00048512706,0.00070359494,0.48309,0.0055903406,0.2175808,0.28193146,0.00014192385],"about_ca_topic_score_codex":0.0061634663,"about_ca_topic_score_gemma":0.0061279833,"teacher_disagreement_score":0.0061634663,"about_ca_system_score_codex":0.0011219254,"about_ca_system_score_gemma":0.0020381159,"threshold_uncertainty_score":0.012255192},"labels":[],"label_agreement":null},{"id":"W4306317452","doi":"10.1145/3511808.3557209","title":"Named Entity-based Question-Answering Pair Generator","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 31st ACM International Conference on Information &amp; Knowledge Management","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Pipeline (software); Paragraph; Question answering; Generator (circuit theory); Task (project management); Context (archaeology); Abstraction; Simple (philosophy); Text generation; Natural language processing; Argument (complex analysis); Artificial intelligence; Programming language; Engineering; World Wide Web","score_opus":0.045095399802941936,"score_gpt":0.2833785768154452,"score_spread":0.23828317701250326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4306317452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046177506,0.00009459643,0.97003525,0.0002953246,0.00012809885,0.00058557256,0.0014460478,0.020701405,0.002095946],"genre_scores_gemma":[0.13940625,0.000109010194,0.8399238,0.0004254933,0.00011928515,0.001443091,0.009396971,0.0029019401,0.006274046],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.995666,0.002139225,0.00027462232,0.000994851,0.00071974756,0.0002054646],"domain_scores_gemma":[0.99091303,0.005523609,0.0002582547,0.0016095327,0.0014509341,0.00024467692],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056491415,0.0018355218,0.0010365221,0.0019780037,0.0007890068,0.001408177,0.0031809756,0.0021718463,0.023262452],"category_scores_gemma":[0.017096499,0.00085079466,0.0016947848,0.0011835333,0.000887506,0.0035988207,0.0040639704,0.0018607005,0.009585091],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021884697,0.0007497321,0.005234408,0.0019978895,0.0002931121,0.0024628022,0.003276229,0.036230434,0.070249565,0.12178329,0.11549504,0.6400391],"study_design_scores_gemma":[0.00044060487,0.00051807234,0.0013973179,0.00010769587,0.00017569629,0.0013040059,0.0006870131,0.6634192,0.10417069,0.121833384,0.10577483,0.00017149589],"about_ca_topic_score_codex":0.0010719105,"about_ca_topic_score_gemma":0.0009091084,"teacher_disagreement_score":0.023262452,"about_ca_system_score_codex":0.0008689616,"about_ca_system_score_gemma":0.0010886509,"threshold_uncertainty_score":0.0778206},"labels":[],"label_agreement":null},{"id":"W4306317500","doi":"10.1145/3511808.3557588","title":"Early Stage Sparse Retrieval with Entity Linking","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 31st ACM International Conference on Information &amp; Knowledge Management","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Question answering; Boosting (machine learning); Artificial intelligence; Task (project management); Ranking (information retrieval); Natural language processing","score_opus":0.05496141773196674,"score_gpt":0.27540141818841524,"score_spread":0.2204400004564485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4306317500","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06778096,0.0012324497,0.8983613,0.00030962395,0.00012697153,0.0006124886,0.0014136137,0.024679366,0.005483081],"genre_scores_gemma":[0.33934486,0.0007476377,0.63275987,0.00051758345,0.00020750758,0.00029703867,0.008613119,0.0008165161,0.01669584],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998723,0.00024135705,0.00012339182,0.00030316267,0.00046066788,0.00014850643],"domain_scores_gemma":[0.99732876,0.00067173876,0.00016228382,0.0013076103,0.00045381952,0.000075827185],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017151574,0.0012337473,0.0016003005,0.002927542,0.0007057323,0.0017922898,0.0023460882,0.0013689675,0.006752098],"category_scores_gemma":[0.005650797,0.00052819814,0.0011263834,0.0035131073,0.0007875248,0.006290796,0.0032371846,0.0011957581,0.006641953],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070087705,0.0007517256,0.0035675678,0.0006168268,0.00020132217,0.00042793175,0.000474772,0.048112407,0.056503113,0.00998266,0.028508076,0.8501527],"study_design_scores_gemma":[0.00022775614,0.000895071,0.002937778,0.000056583547,0.0002338926,0.0013181872,0.0003458598,0.837341,0.09633876,0.028525645,0.031638116,0.0001413078],"about_ca_topic_score_codex":0.0043069515,"about_ca_topic_score_gemma":0.008328473,"teacher_disagreement_score":0.006752098,"about_ca_system_score_codex":0.0005510416,"about_ca_system_score_gemma":0.001233929,"threshold_uncertainty_score":0.022588015},"labels":[],"label_agreement":null},{"id":"W4306317650","doi":"10.1145/3511808.3557719","title":"Unsupervised Question Clarity Prediction through Retrieved Item Coherency","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 31st ACM International Conference on Information &amp; Knowledge Management","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Ambiguity; Computer science; CLARITY; Generalization; Context (archaeology); Artificial intelligence; Ask price; Graph; Machine learning; Similarity (geometry); Open domain; Information retrieval; Question answering; Natural language processing; Theoretical computer science; Mathematics","score_opus":0.06448416126311513,"score_gpt":0.29336109989716924,"score_spread":0.22887693863405412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4306317650","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7192603,0.0058351317,0.250825,0.0010498957,0.00013567734,0.0005234069,0.0075628245,0.007918127,0.006889749],"genre_scores_gemma":[0.91354585,0.0004828729,0.06912764,0.00017207736,0.0001128879,0.00017329487,0.013817697,0.00015105984,0.0024166715],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983342,0.0005522562,0.00013204201,0.00054985663,0.0002954112,0.00013633251],"domain_scores_gemma":[0.9933555,0.0040762653,0.0005236068,0.00051145715,0.0012908246,0.00024232373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021226802,0.0009489417,0.00089187466,0.004417122,0.0005379682,0.001197417,0.0013580763,0.0013820304,0.0012075589],"category_scores_gemma":[0.0108595835,0.0002915529,0.0008947094,0.0021388433,0.0005223381,0.0029869117,0.001369463,0.0012718872,0.0010894368],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023712816,0.0011414626,0.13933122,0.0016809942,0.000816518,0.00064649637,0.003351554,0.06821698,0.06273989,0.0047639366,0.028249357,0.6866903],"study_design_scores_gemma":[0.0000766503,0.00031228602,0.047783416,0.00006496491,0.00027458038,0.00032834426,0.0008324197,0.9149299,0.018114436,0.008543884,0.008664644,0.00007442768],"about_ca_topic_score_codex":0.008384282,"about_ca_topic_score_gemma":0.016468152,"teacher_disagreement_score":0.008384282,"about_ca_system_score_codex":0.00072848937,"about_ca_system_score_gemma":0.0009421289,"threshold_uncertainty_score":0.016670942},"labels":[],"label_agreement":null},{"id":"W4307895180","doi":"10.48550/arxiv.2205.09393","title":"Two-Step Question Retrieval for Open-Domain QA","year":2022,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Search engine indexing; Inference; Computer science; Labrador Retriever; Information retrieval; Pipeline (software); Domain (mathematical analysis); Squid; Artificial intelligence; Mathematics; Programming language; Medicine; Biology; Fishery","score_opus":0.09835796420991373,"score_gpt":0.23386526790810672,"score_spread":0.13550730369819297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4307895180","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022062365,0.0009434015,0.9639532,0.00070406776,0.00008172502,0.0002927185,0.0006084961,0.009037466,0.0023165704],"genre_scores_gemma":[0.50212836,0.0005451936,0.48616016,0.0006148547,0.00018894857,0.00041202342,0.0033390485,0.00040270315,0.006208631],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986356,0.0005653143,0.00007940084,0.00039129375,0.0002171256,0.00011127522],"domain_scores_gemma":[0.9948443,0.00283721,0.00017141056,0.0013078782,0.0006351819,0.00020406568],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033194718,0.0008214837,0.0010548156,0.0012368354,0.00062822993,0.0011377129,0.0025366575,0.0018239686,0.008078161],"category_scores_gemma":[0.0111141335,0.0006167188,0.0012144407,0.0009903305,0.00094292767,0.004812097,0.0027496952,0.0028407155,0.0044047027],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009939384,0.00082549796,0.005413118,0.00090568274,0.00019095505,0.00022585072,0.0008241445,0.1113341,0.022913197,0.034122277,0.033202082,0.7890491],"study_design_scores_gemma":[0.00007152396,0.00015028732,0.0006496636,0.000018162676,0.00003037647,0.00016396982,0.00005755772,0.96347135,0.0070949257,0.023591643,0.0046760645,0.000024472834],"about_ca_topic_score_codex":0.007862336,"about_ca_topic_score_gemma":0.00912227,"teacher_disagreement_score":0.008078161,"about_ca_system_score_codex":0.001155247,"about_ca_system_score_gemma":0.0018643838,"threshold_uncertainty_score":0.02702415},"labels":[],"label_agreement":null},{"id":"W4308641647","doi":"10.1145/3540250.3549145","title":"Are we building on the rock? on the importance of data preprocessing for code summarization","year":2022,"lang":"en","type":"article","venue":"Proceedings of the 30th ACM Joint European Software Engineering Conference and Symposium on the Foundations of Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China; Youth Innovation Promotion Association; Chinese Academy of Sciences; National Science Foundation","keywords":"Automatic summarization; Benchmark (surveying); Computer science; Preprocessor; Code (set theory); Data pre-processing; Data mining; Benchmarking; Machine learning; Task (project management); Artificial intelligence; Engineering","score_opus":0.0714784426833877,"score_gpt":0.2534673684029559,"score_spread":0.1819889257195682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308641647","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15080996,0.026981827,0.6570436,0.12145027,0.0025883422,0.0010448411,0.007429525,0.014406009,0.01824571],"genre_scores_gemma":[0.3780777,0.012331496,0.5662274,0.016562246,0.0016885581,0.0010902529,0.014926626,0.0045912485,0.004504514],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9454184,0.024869028,0.0045575527,0.008594663,0.014884044,0.0016762511],"domain_scores_gemma":[0.6489702,0.2168395,0.016336918,0.060154784,0.053023264,0.004675346],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.049364068,0.0025463316,0.0030900827,0.00801762,0.0044729877,0.0143997045,0.0041854545,0.0048165703,0.0050610546],"category_scores_gemma":[0.2953902,0.0025537058,0.0025438005,0.008039505,0.007752884,0.049092773,0.009436071,0.009671842,0.006035982],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018751341,0.00048888515,0.06908237,0.0041734017,0.0007685868,0.0005602814,0.009928337,0.016401429,0.012562317,0.04241479,0.07850237,0.76324207],"study_design_scores_gemma":[0.00046604674,0.0021309592,0.064343855,0.008643712,0.0011847838,0.0018655182,0.017771192,0.20546916,0.042482015,0.32150483,0.33291036,0.0012275196],"about_ca_topic_score_codex":0.011080068,"about_ca_topic_score_gemma":0.012632756,"teacher_disagreement_score":0.049364068,"about_ca_system_score_codex":0.0025130396,"about_ca_system_score_gemma":0.0077937734,"threshold_uncertainty_score":0.2610653},"labels":[],"label_agreement":null},{"id":"W4309232742","doi":"10.21449/ijate.1124382","title":"Automatic story and item generation for reading comprehension assessments with transformers","year":2022,"lang":"en","type":"article","venue":"International Journal of Assessment Tools in Education","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University of Edmonton; University of Alberta","funders":"University of Alberta","keywords":"Fluency; Reading comprehension; Computer science; Comprehension; Literacy; Reading (process); Mathematics education; Multimedia; Psychology; Pedagogy; Linguistics","score_opus":0.04583102892320276,"score_gpt":0.3728767693014102,"score_spread":0.32704574037820744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309232742","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10672842,0.00012755574,0.8482456,0.0001174348,0.000102521284,0.00085578306,0.002273449,0.03776142,0.0037877972],"genre_scores_gemma":[0.39646822,0.00009370769,0.5943442,0.000051443454,0.000022507826,0.0010468789,0.0045336364,0.0013642772,0.0020751285],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997273,0.0013882989,0.000310145,0.00052363455,0.00042451144,0.000080288555],"domain_scores_gemma":[0.98727137,0.008587195,0.0005099711,0.0013653706,0.0020204287,0.0002457183],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002884775,0.0011876706,0.000641768,0.0020629673,0.00032147823,0.0017987638,0.0012864015,0.00079552067,0.0085713845],"category_scores_gemma":[0.025113447,0.0004978523,0.0008321128,0.0012328271,0.0003329142,0.0028965934,0.0016909826,0.0008385849,0.0038193986],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010455959,0.0005025618,0.015986908,0.0006811723,0.0001447017,0.0005826881,0.0030845448,0.02409746,0.040167272,0.005469442,0.011679709,0.89655787],"study_design_scores_gemma":[0.00028984685,0.00065198634,0.012898557,0.00011939548,0.00013098729,0.0007524031,0.001241709,0.87789565,0.07065073,0.011852963,0.023383224,0.00013247461],"about_ca_topic_score_codex":0.0013592726,"about_ca_topic_score_gemma":0.0018380581,"teacher_disagreement_score":0.0085713845,"about_ca_system_score_codex":0.00052075397,"about_ca_system_score_gemma":0.00071688317,"threshold_uncertainty_score":0.028674126},"labels":[],"label_agreement":null},{"id":"W4309375705","doi":"10.1080/15305058.2022.2070755","title":"Generating reading comprehension items using automated processes","year":2022,"lang":"en","type":"article","venue":"International Journal of Testing","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Reading comprehension; Blank; Comprehension; Computer science; Natural language processing; Test (biology); Reading (process); Artificial intelligence; Salient; Process (computing); Psychology; Linguistics","score_opus":0.08765853027184589,"score_gpt":0.32709880228236776,"score_spread":0.23944027201052187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309375705","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03576569,0.00006690312,0.94650364,0.00010586727,0.000039103186,0.0010965605,0.000802841,0.013969071,0.0016503228],"genre_scores_gemma":[0.121783085,0.00006730525,0.86923736,0.00007398934,0.00004312238,0.0022137314,0.0034034154,0.0009616341,0.0022163412],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9938393,0.002801723,0.0004933787,0.0018287277,0.0009000025,0.00013683349],"domain_scores_gemma":[0.9609355,0.029033665,0.0013088643,0.0037262721,0.004781394,0.0002143193],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005174801,0.0021493,0.0011580989,0.003313177,0.0008187149,0.0018939304,0.0019018549,0.001229293,0.009178868],"category_scores_gemma":[0.034333766,0.0008219539,0.0014698132,0.002132979,0.0008327854,0.0021395986,0.002091763,0.0017315635,0.006756426],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034772264,0.0006516719,0.008139097,0.0005188125,0.000120247096,0.00029167414,0.0026713668,0.011731364,0.034311622,0.003541599,0.0054219733,0.93225294],"study_design_scores_gemma":[0.0003609036,0.0011690138,0.02351566,0.0001588897,0.00025694893,0.00082626747,0.0017926564,0.79158497,0.12684284,0.027960006,0.025300218,0.0002316466],"about_ca_topic_score_codex":0.001508119,"about_ca_topic_score_gemma":0.0019222962,"teacher_disagreement_score":0.009178868,"about_ca_system_score_codex":0.0006868548,"about_ca_system_score_gemma":0.0012726821,"threshold_uncertainty_score":0.030706406},"labels":[],"label_agreement":null},{"id":"W4309663019","doi":"10.1126/science.ade9097","title":"Human-level play in the game of<i>Diplomacy</i>by combining language models with strategic reasoning","year":2022,"lang":"en","type":"article","venue":"Science","topic":"Topic Modeling","field":"Computer Science","cited_by":201,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Cicero; Negotiation; Diplomacy; Computer science; Competition (biology); Reinforcement learning; Artificial intelligence; League; Natural language; Political science; Politics; Law; History; Ecology","score_opus":0.051099470607239225,"score_gpt":0.2899594067233681,"score_spread":0.23885993611612888,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309663019","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50914377,0.0003638365,0.42848492,0.002164998,0.00010828335,0.00045202518,0.00031667508,0.0024159693,0.056549434],"genre_scores_gemma":[0.895644,0.00012250946,0.099444404,0.000163013,0.000016936456,0.00011347965,0.00025661103,0.00006688955,0.0041720797],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99912673,0.00048401096,0.000034916917,0.00014567528,0.00012477,0.00008384605],"domain_scores_gemma":[0.99781054,0.00135318,0.00021251712,0.00020507038,0.00017351378,0.00024508318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020235386,0.0009239783,0.00032935358,0.0004117798,0.0005520736,0.002604939,0.0009987416,0.0006266536,0.0025431097],"category_scores_gemma":[0.0052165114,0.00029062707,0.00044519972,0.00014344919,0.001419059,0.0017957593,0.0013945292,0.0011150643,0.000628232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014175582,0.0018519114,0.06395772,0.00094904023,0.0004977329,0.0006135692,0.012911878,0.34457776,0.047468487,0.20196213,0.02061604,0.3031761],"study_design_scores_gemma":[0.00010339051,0.00071271247,0.0076861926,0.00010223666,0.00007450008,0.00021309615,0.001956231,0.87989414,0.010654361,0.06081417,0.037690695,0.00009817799],"about_ca_topic_score_codex":0.01122816,"about_ca_topic_score_gemma":0.014680199,"teacher_disagreement_score":0.01122816,"about_ca_system_score_codex":0.0012579832,"about_ca_system_score_gemma":0.0017651317,"threshold_uncertainty_score":0.022325635},"labels":[],"label_agreement":null},{"id":"W4310064233","doi":"10.1016/j.neunet.2022.11.028","title":"SGORNN: Combining scalar gates and orthogonal constraints in recurrent networks","year":2022,"lang":"en","type":"article","venue":"Neural Networks","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Horizon 2020; Natural Sciences and Engineering Research Council of Canada; Narodowe Centrum Nauki","keywords":"Recurrent neural network; Overfitting; Scalar (mathematics); Computer science; Treebank; Probabilistic logic; Algorithm; Artificial intelligence; Context (archaeology); Deep learning; Backpropagation; Artificial neural network; Pattern recognition (psychology); Mathematics; Annotation","score_opus":0.017934032237355308,"score_gpt":0.2324914549739667,"score_spread":0.21455742273661138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4310064233","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011957191,0.00024595778,0.97823036,0.00016783278,0.00012941683,0.00008095718,0.0004908431,0.0071862056,0.001511315],"genre_scores_gemma":[0.2959864,0.00038349285,0.6909811,0.00042154492,0.00016580863,0.00033818602,0.0031976348,0.0018481851,0.0066776443],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99923325,0.0002900048,0.00007201565,0.00017784348,0.00014417939,0.000082751285],"domain_scores_gemma":[0.9987249,0.0006364068,0.00006997561,0.000245827,0.00024223747,0.000080550344],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022941881,0.0011322211,0.0011295477,0.00060652033,0.00033423427,0.0012652713,0.0017361118,0.0009815665,0.005082559],"category_scores_gemma":[0.0051976065,0.0006325268,0.0008137562,0.00086599955,0.00047241047,0.002767409,0.0018641878,0.0019036861,0.0020856843],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000552374,0.00029642184,0.0011582298,0.00032309367,0.00024589893,0.00016619528,0.00013904955,0.26474366,0.0149996,0.028440244,0.019214265,0.6697209],"study_design_scores_gemma":[0.00002351235,0.0000354859,0.00007110415,0.000008745664,0.000018570236,0.000012141178,0.000006540792,0.98715436,0.002589303,0.008651432,0.0014197604,0.0000089924315],"about_ca_topic_score_codex":0.0052991044,"about_ca_topic_score_gemma":0.013181485,"teacher_disagreement_score":0.0052991044,"about_ca_system_score_codex":0.000524977,"about_ca_system_score_gemma":0.0012746498,"threshold_uncertainty_score":0.017002821},"labels":[],"label_agreement":null},{"id":"W4310368227","doi":"10.1101/2022.11.28.22282767","title":"Natural Language Processing for Clinical Laboratory Data Repository Systems: Implementation and Evaluation for Respiratory Viruses","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sunnybrook Hospital; Sinai Health System; Institute for Clinical Evaluative Sciences; Vector Institute; Public Health Ontario; University Health Network; University of Toronto","funders":"Vector Institute; Canadian Institutes of Health Research; Hospital for Sick Children","keywords":"Computer science; Artificial intelligence; Natural language processing; Generalizability theory; Parsing; Machine learning; Classifier (UML); F1 score; Information extraction","score_opus":0.24070224106144006,"score_gpt":0.49002221284359226,"score_spread":0.2493199717821522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4310368227","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75489026,0.0013961381,0.14369904,0.0018703685,0.00041370836,0.003332363,0.006193922,0.08307398,0.0051301965],"genre_scores_gemma":[0.7083008,0.00068054214,0.27023456,0.00055316894,0.00005704022,0.0012505744,0.014518581,0.0009340377,0.0034706553],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969766,0.0010304231,0.00036228457,0.00086545455,0.00058382814,0.00018146327],"domain_scores_gemma":[0.99204004,0.004794008,0.00032581473,0.0007095373,0.0017231803,0.0004074623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060903504,0.0011476289,0.0006161483,0.0010774184,0.00070848176,0.0013676416,0.0037095994,0.0015506381,0.0032427437],"category_scores_gemma":[0.015106965,0.00056798063,0.00076184765,0.0010020675,0.00070052216,0.0026827075,0.0017113767,0.0020725676,0.0019275537],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003545332,0.0057342895,0.030676309,0.0031712772,0.0009161049,0.0020960572,0.0024499854,0.17758782,0.040606946,0.0021443258,0.052266877,0.67880476],"study_design_scores_gemma":[0.000384839,0.00067633233,0.0057323594,0.0000967606,0.00012426973,0.00025630853,0.00046856792,0.9571471,0.027077219,0.0010208413,0.0069435113,0.00007179715],"about_ca_topic_score_codex":0.025077432,"about_ca_topic_score_gemma":0.019414976,"teacher_disagreement_score":0.025077432,"about_ca_system_score_codex":0.0021670042,"about_ca_system_score_gemma":0.0029598542,"threshold_uncertainty_score":0.04986292},"labels":[],"label_agreement":null},{"id":"W4310990713","doi":"10.1111/exsy.13183","title":"Temporal positional lexicon expansion for federated learning based on hyperpatism detection","year":2022,"lang":"en","type":"article","venue":"Expert Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brandon University","funders":"","keywords":"Computer science; Artificial intelligence; Focus (optics); Deep learning; Lexicon; Artificial neural network; Machine learning; Supervised learning; Social media; Natural language processing","score_opus":0.025501309991139105,"score_gpt":0.25487780854000847,"score_spread":0.22937649854886935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4310990713","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.375437,0.0005241389,0.6072803,0.00046520002,0.0001211359,0.00017668826,0.00065883755,0.011328716,0.0040079523],"genre_scores_gemma":[0.8935319,0.00006477146,0.10125699,0.00012633437,0.000046697744,0.00010167849,0.0013329553,0.00012700034,0.0034117114],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99943334,0.0001593677,0.00004187297,0.0002080111,0.000090768866,0.00006669072],"domain_scores_gemma":[0.99797875,0.0010060188,0.00016372958,0.00028019233,0.00047725684,0.00009411879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010551523,0.00077960116,0.0008030763,0.0015810702,0.0005498226,0.00082696456,0.001423555,0.000852886,0.0018897733],"category_scores_gemma":[0.003280475,0.0003852869,0.0007612311,0.0008999736,0.00050356856,0.0018980652,0.0011920751,0.0009928718,0.0007763111],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004327631,0.00080604205,0.007206974,0.00010794837,0.00009112575,0.000324742,0.00031192484,0.2268422,0.018617531,0.0044196155,0.007006659,0.7338324],"study_design_scores_gemma":[0.0000064652604,0.000021436019,0.00027184418,0.0000021540725,0.0000047885082,0.00001537943,0.000013819994,0.9966515,0.001729838,0.0010418049,0.00023681766,0.0000041208327],"about_ca_topic_score_codex":0.007324193,"about_ca_topic_score_gemma":0.0089864945,"teacher_disagreement_score":0.007324193,"about_ca_system_score_codex":0.0008809504,"about_ca_system_score_gemma":0.0010936086,"threshold_uncertainty_score":0.014563143},"labels":[],"label_agreement":null},{"id":"W4311617908","doi":"10.1101/2022.11.30.22282946","title":"Discovering Social Determinants of Health from Case Reports using Natural Language Processing: Algorithmic Development and Validation","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"Institute of Health Services and Policy Research; Canadian Institutes of Health Research","keywords":"Computer science; Benchmark (surveying); Natural language processing; Artificial intelligence; Social media; Social determinants of health; Annotation; Key (lock); Information extraction; Set (abstract data type); Information retrieval; Health care; Data science; Machine learning; World Wide Web; Political science","score_opus":0.055179389816222425,"score_gpt":0.335244282657424,"score_spread":0.28006489284120156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4311617908","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16858365,0.0014278523,0.81646585,0.0015746474,0.00010818951,0.0019530866,0.0034220172,0.0044629006,0.0020018106],"genre_scores_gemma":[0.31641293,0.00039675774,0.673245,0.00019710246,0.00011534025,0.0012916265,0.0077490634,0.00011060622,0.000481563],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.991757,0.004875994,0.00088374555,0.001506365,0.0008112029,0.00016573927],"domain_scores_gemma":[0.918048,0.074048944,0.0018162039,0.0028771495,0.0029272223,0.00028255326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015243094,0.001094777,0.0006454711,0.005391206,0.000856857,0.0019853322,0.0023260035,0.0016098967,0.0018830524],"category_scores_gemma":[0.041960187,0.0005029545,0.0014629464,0.0020791416,0.0013570032,0.0017914412,0.0019925295,0.0016150214,0.00076791964],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005735337,0.0013692349,0.055159003,0.001659904,0.00056018075,0.0015045332,0.0017969988,0.22362705,0.0064277826,0.008263349,0.012640473,0.686418],"study_design_scores_gemma":[0.00008127803,0.00008188894,0.00625639,0.00013212682,0.00008216971,0.00034587315,0.00038402708,0.97584265,0.0032117984,0.0108962385,0.002661613,0.00002395719],"about_ca_topic_score_codex":0.0051894863,"about_ca_topic_score_gemma":0.0060639963,"teacher_disagreement_score":0.015243094,"about_ca_system_score_codex":0.0011307617,"about_ca_system_score_gemma":0.0023126968,"threshold_uncertainty_score":0.08061415},"labels":[],"label_agreement":null},{"id":"W4311706555","doi":"10.1002/aaai.12068","title":"Search and learning for unsupervised text generation","year":2022,"lang":"en","type":"article","venue":"AI Magazine","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"DeepMind; Alberta Machine Intelligence Institute; Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Artificial intelligence; Heuristic; Task (project management); Sentence; Annotation; Machine learning; Function (biology); Component (thermodynamics); Unsupervised learning; Natural language processing; Resource (disambiguation)","score_opus":0.039183752693572334,"score_gpt":0.2721839684174112,"score_spread":0.2330002157238389,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4311706555","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011601858,0.00037459686,0.9841339,0.0005035927,0.000035054978,0.000047284047,0.00014673201,0.00066938257,0.0024875677],"genre_scores_gemma":[0.47637087,0.0006139613,0.5103254,0.00040022994,0.00021314697,0.00040269655,0.0011768983,0.00041628443,0.010080478],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99932027,0.00029598092,0.00003898583,0.00016820921,0.00013309631,0.00004352926],"domain_scores_gemma":[0.99709463,0.0020909098,0.00019425439,0.00027937867,0.00027267638,0.000068229776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013460135,0.000353053,0.0006473921,0.0010338027,0.0005100503,0.00097050617,0.0009406055,0.0008150227,0.0038809355],"category_scores_gemma":[0.005848308,0.00032874942,0.00079854194,0.001080765,0.0009578404,0.0017757931,0.0011041814,0.0012670549,0.00093077845],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000106129584,0.00010806697,0.0013175791,0.00024056646,0.000063134234,0.000110439745,0.00019721134,0.5453834,0.005031321,0.175144,0.013367876,0.2589303],"study_design_scores_gemma":[0.000006852437,0.000008466731,0.00006829176,0.000006899348,0.000002619484,0.000013907457,0.000007659296,0.9580219,0.00061563635,0.04001963,0.0012242759,0.0000037794243],"about_ca_topic_score_codex":0.0016262563,"about_ca_topic_score_gemma":0.002586413,"teacher_disagreement_score":0.0038809355,"about_ca_system_score_codex":0.0010171521,"about_ca_system_score_gemma":0.00090300443,"threshold_uncertainty_score":0.012983084},"labels":[],"label_agreement":null},{"id":"W4312001626","doi":"10.5121/csit.2022.122210","title":"Word Embedding Interpretation using Co-Clustering","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Carleton University","funders":"","keywords":"Word embedding; Word (group theory); Computer science; Embedding; Natural language processing; Artificial intelligence; Cluster analysis; Interpretability; Representation (politics); Simple (philosophy); Linguistics","score_opus":0.03765871830058227,"score_gpt":0.31166386780476596,"score_spread":0.2740051495041837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312001626","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019392857,0.00064407755,0.97404534,0.00015077517,0.00013990726,0.00012534426,0.00047619443,0.0029826737,0.002042814],"genre_scores_gemma":[0.25766093,0.0007717416,0.7277486,0.000125193,0.00020114004,0.0002659879,0.004132151,0.0009752453,0.00811907],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99834347,0.00032984547,0.00013048471,0.000648156,0.0004086666,0.00013936014],"domain_scores_gemma":[0.9972863,0.000674601,0.00021727377,0.00056975585,0.001170063,0.00008187819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007174754,0.0019582943,0.0010608439,0.006047579,0.0009450408,0.0018197143,0.0012514592,0.001380686,0.0033980946],"category_scores_gemma":[0.004883487,0.00042479107,0.0013835765,0.005129203,0.00075907103,0.0023602645,0.0020681648,0.0014753971,0.0035589172],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037408547,0.00020364439,0.0036227538,0.00053935393,0.00029466645,0.00043034047,0.001291942,0.0377406,0.04130652,0.021588756,0.014989436,0.877618],"study_design_scores_gemma":[0.000042061827,0.00014512778,0.0026640682,0.00007627223,0.00016625023,0.000580993,0.00071527547,0.8796891,0.033391666,0.05980724,0.022613412,0.000108473876],"about_ca_topic_score_codex":0.0039865826,"about_ca_topic_score_gemma":0.005855057,"teacher_disagreement_score":0.006047579,"about_ca_system_score_codex":0.00059402006,"about_ca_system_score_gemma":0.0011197323,"threshold_uncertainty_score":0.011367738},"labels":[],"label_agreement":null},{"id":"W4312116152","doi":"10.2196/preprints.45268","title":"Leveraging Knowledge Graphs and Natural Language Processing for Automated Web Resource Labeling and Knowledge Mobilization in Neurodevelopmental Disorders: Development and Usability Study (Preprint)","year":2022,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Women and Children’s Health Research Institute; University of Alberta","funders":"","keywords":"Computer science; Terminology; Usability; World Wide Web; Artificial intelligence; Data science; Knowledge management; Information retrieval; Natural language processing","score_opus":0.02299959624589169,"score_gpt":0.2972719642577111,"score_spread":0.27427236801181937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312116152","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69985056,0.0010147544,0.25993675,0.0013669359,0.00013953327,0.006631491,0.007993252,0.017098118,0.0059687486],"genre_scores_gemma":[0.39913002,0.00052489777,0.58715713,0.0003388339,0.000022123884,0.002667525,0.008051101,0.00067818473,0.0014302152],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99305433,0.0044171955,0.0007040506,0.0009460813,0.000730785,0.00014763452],"domain_scores_gemma":[0.91524136,0.07691569,0.0014074879,0.0025972854,0.003429949,0.00040828958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0124997245,0.00085953675,0.00057524216,0.0039418107,0.0007269942,0.002637331,0.0012716377,0.00082482985,0.002146554],"category_scores_gemma":[0.04258544,0.00053299183,0.0013889273,0.0022924077,0.00076695683,0.004791966,0.0028327142,0.0012261384,0.00070721656],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010107559,0.0028538404,0.03038873,0.006598934,0.00064410333,0.0010358898,0.030489478,0.0151059255,0.029741855,0.004581472,0.01952736,0.8580216],"study_design_scores_gemma":[0.0008025313,0.004619897,0.106088266,0.0030574496,0.0014635198,0.001969167,0.03969354,0.6144916,0.07510022,0.024810974,0.12699993,0.0009029602],"about_ca_topic_score_codex":0.008625401,"about_ca_topic_score_gemma":0.011845237,"teacher_disagreement_score":0.0124997245,"about_ca_system_score_codex":0.0013085203,"about_ca_system_score_gemma":0.0017843397,"threshold_uncertainty_score":0.066105604},"labels":[],"label_agreement":null},{"id":"W4312125493","doi":"10.3390/v14122761","title":"Clinical Application of Detecting COVID-19 Risks: A Natural Language Processing Approach","year":2022,"lang":"en","type":"article","venue":"Viruses","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto; Toronto Metropolitan University","funders":"","keywords":"Novelty; Computer science; Machine learning; Artificial intelligence; Artificial neural network; Pipeline (software); Pandemic; Transformer; F1 score; Task (project management); Named-entity recognition; Coronavirus disease 2019 (COVID-19); Natural language processing; Infectious disease (medical specialty); Disease; Medicine; Psychology; Engineering","score_opus":0.11799924662285974,"score_gpt":0.404890190535899,"score_spread":0.28689094391303926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312125493","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16429454,0.0065571847,0.7823772,0.009431002,0.00069457263,0.0016320241,0.021774873,0.005130467,0.008108162],"genre_scores_gemma":[0.5864234,0.0020101957,0.38949457,0.0013290789,0.00078601827,0.00065055816,0.01649764,0.00012431806,0.002684281],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99789006,0.00068393245,0.00034413888,0.00061412086,0.00033709762,0.000130657],"domain_scores_gemma":[0.9948638,0.0032869896,0.0006514195,0.0003136982,0.000710638,0.00017344789],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023784235,0.0009320679,0.00065099425,0.0049401037,0.0004901198,0.0013602918,0.00089748856,0.001310823,0.001543251],"category_scores_gemma":[0.006996361,0.0002661856,0.0012353837,0.0020491164,0.00045368125,0.0014675133,0.0010579197,0.0012745891,0.0010794532],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000892185,0.0009387597,0.118454136,0.0019259896,0.00041736467,0.0037699125,0.0013819895,0.023169823,0.042403597,0.010452094,0.037307188,0.75888693],"study_design_scores_gemma":[0.00013206276,0.00054240256,0.07245379,0.00040945,0.0004632412,0.006421732,0.0018309491,0.7953373,0.025226228,0.047214568,0.049761713,0.0002066419],"about_ca_topic_score_codex":0.0051347907,"about_ca_topic_score_gemma":0.0065856497,"teacher_disagreement_score":0.0051347907,"about_ca_system_score_codex":0.00070370146,"about_ca_system_score_gemma":0.0016118384,"threshold_uncertainty_score":0.012578428},"labels":[],"label_agreement":null},{"id":"W4312217620","doi":"10.18280/ria.360516","title":"Improving Extractive Text Summarization Performance Using Enhanced Feature Based RBM Method","year":2022,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Discriminative model; Artificial intelligence; Feature (linguistics); Feature selection; Restricted Boltzmann machine; Set (abstract data type); Word (group theory); Natural language processing; Sentence; Feature extraction; Topic model; Multi-document summarization; Process (computing); Information retrieval; Artificial neural network; Pattern recognition (psychology)","score_opus":0.04148329245521411,"score_gpt":0.2896778044399188,"score_spread":0.2481945119847047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312217620","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06326697,0.002887624,0.92138934,0.0003485254,0.00024654384,0.00017371788,0.0005780797,0.008756666,0.0023524403],"genre_scores_gemma":[0.39653963,0.0011764338,0.5862473,0.00021584093,0.00034575246,0.00036128375,0.0031309575,0.00041709526,0.011565689],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993499,0.00012840277,0.000078580015,0.00015704593,0.00022033561,0.00006568436],"domain_scores_gemma":[0.9991522,0.00022591693,0.00012525721,0.00009184985,0.00036839355,0.000036350077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00072227593,0.0009568768,0.0011904071,0.0015779141,0.00041643414,0.000694733,0.0009185086,0.0007011099,0.001984179],"category_scores_gemma":[0.0017666034,0.00023927358,0.0009106086,0.0010123608,0.00022304135,0.001218484,0.00041592176,0.0006943816,0.0017410086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042006015,0.00023332381,0.0015156377,0.00037400602,0.00012794799,0.00030802458,0.0002140385,0.05315784,0.11631708,0.0015328018,0.008598448,0.8172009],"study_design_scores_gemma":[0.00007355743,0.00037158842,0.0029123742,0.000023471162,0.00013270398,0.00029845216,0.0001096841,0.91748583,0.06873268,0.0017084267,0.008106472,0.000044729262],"about_ca_topic_score_codex":0.0028709238,"about_ca_topic_score_gemma":0.0038441848,"teacher_disagreement_score":0.0028709238,"about_ca_system_score_codex":0.0004553476,"about_ca_system_score_gemma":0.0007108377,"threshold_uncertainty_score":0.006637752},"labels":[],"label_agreement":null},{"id":"W4312386790","doi":"10.1145/3524610.3527892","title":"On the cross-modal transfer from natural language to code through adapter modules","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Adapter (computing); Computer science; Source code; Natural language; Software; Code (set theory); Programming language; Modal; Reverse engineering; Artificial intelligence; Natural language processing; Computer hardware","score_opus":0.024742572356654505,"score_gpt":0.27140193072297136,"score_spread":0.24665935836631686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312386790","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06774984,0.0006751152,0.91368145,0.0010204167,0.00010933529,0.00016139088,0.0002763218,0.0075253802,0.008800788],"genre_scores_gemma":[0.70639783,0.0012312395,0.27400398,0.001177972,0.00013719144,0.00045842308,0.0016305128,0.00189934,0.013063455],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983381,0.0007966396,0.00006314321,0.00046438293,0.00019658166,0.00014106777],"domain_scores_gemma":[0.99381644,0.003540433,0.00020478353,0.0017598192,0.0005265728,0.00015187917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036427577,0.0011698595,0.0005261841,0.00070713775,0.00045815218,0.0016604506,0.0018710982,0.0012701342,0.006954365],"category_scores_gemma":[0.016376082,0.0006132883,0.0010262282,0.0006406173,0.0015379824,0.006367968,0.005428671,0.002768986,0.0031445364],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065582094,0.0004361036,0.0040285722,0.00039836532,0.00028247465,0.0004259675,0.002049597,0.17181088,0.031951483,0.055147514,0.010149953,0.72266334],"study_design_scores_gemma":[0.00002910717,0.0002247839,0.0014816984,0.00008618948,0.00008196773,0.00016422755,0.00028633574,0.91213447,0.013771282,0.063994296,0.0077056815,0.000039976585],"about_ca_topic_score_codex":0.004198018,"about_ca_topic_score_gemma":0.003685196,"teacher_disagreement_score":0.006954365,"about_ca_system_score_codex":0.00088018976,"about_ca_system_score_gemma":0.00081830233,"threshold_uncertainty_score":0.023264706},"labels":[],"label_agreement":null},{"id":"W4312515853","doi":"10.1109/iisa56318.2022.9904390","title":"Question Answering Using Semantic Query Graphs: A Replication Study","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Question answering; Artificial intelligence; RDF; Information retrieval; Natural language processing; Domain (mathematical analysis); Graph; Deep learning; Knowledge graph; Semantic Web; Theoretical computer science","score_opus":0.043269588013825876,"score_gpt":0.30185436646591696,"score_spread":0.2585847784520911,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312515853","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75289804,0.012338983,0.15729874,0.008933481,0.0015035558,0.0049098507,0.0113429325,0.012608511,0.038165838],"genre_scores_gemma":[0.87574935,0.0016380906,0.095768124,0.0023668613,0.00053094974,0.0013905305,0.0153033575,0.0010281599,0.006224453],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97864485,0.014101713,0.0010679407,0.0030000508,0.0027489713,0.0004364665],"domain_scores_gemma":[0.8971064,0.056492172,0.0017826896,0.031354114,0.012165875,0.0010987922],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.021091107,0.0016917911,0.0015190222,0.0033910014,0.0011714023,0.0021853487,0.0029841622,0.0023705876,0.005432954],"category_scores_gemma":[0.08718324,0.0005225166,0.0027068418,0.002092436,0.0019072709,0.008673845,0.003050174,0.002883948,0.0030633933],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005331437,0.010141579,0.07697681,0.0052354746,0.0021677862,0.00085084146,0.0068006096,0.03809,0.026871499,0.019869946,0.07979756,0.7278665],"study_design_scores_gemma":[0.0033659185,0.008457299,0.099791266,0.0016797852,0.003097952,0.0027843972,0.0074881557,0.50785005,0.05337635,0.065175965,0.24620534,0.00072753255],"about_ca_topic_score_codex":0.018100565,"about_ca_topic_score_gemma":0.006981225,"teacher_disagreement_score":0.9789089,"about_ca_system_score_codex":0.0021941566,"about_ca_system_score_gemma":0.0022547278,"threshold_uncertainty_score":0.11154175},"labels":[],"label_agreement":null},{"id":"W4312516176","doi":"10.1162/tacl_a_00511","title":"Causal Inference in Natural Language Processing: Estimation, Prediction, Interpretation and Beyond","year":2022,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":194,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Columbia College","funders":"","keywords":"Causal inference; Computer science; Interpretability; Inference; Artificial intelligence; Causality (physics); Natural language processing; Robustness (evolution); Machine learning; Interpretation (philosophy); Data science; Econometrics","score_opus":0.008751666553274558,"score_gpt":0.26541263141212373,"score_spread":0.25666096485884915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312516176","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006548898,0.008914301,0.9654197,0.014036601,0.00024237428,0.00009624548,0.0004610413,0.0004408295,0.0038400418],"genre_scores_gemma":[0.6168602,0.01481231,0.3547273,0.004995914,0.0036075239,0.0006780988,0.0013720113,0.0004345402,0.0025121071],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9587908,0.033040702,0.0013193854,0.0035924145,0.0028366959,0.000419992],"domain_scores_gemma":[0.6363848,0.3413191,0.0070098923,0.009663693,0.0048919134,0.00073055405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.052185364,0.001321264,0.0025606796,0.0055780527,0.0018762801,0.009268486,0.0036120897,0.0031571547,0.0057395264],"category_scores_gemma":[0.22255951,0.0013560724,0.0021597391,0.005388177,0.011311748,0.013101841,0.005131652,0.007863645,0.00091602566],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012695204,0.000083097504,0.007146084,0.0009523456,0.00046653298,0.00024919538,0.0010441217,0.038596474,0.0003530164,0.8501629,0.0061941594,0.09462497],"study_design_scores_gemma":[0.0000133696785,0.00000948745,0.00051978166,0.00022890502,0.000034480116,0.000041199033,0.00006605686,0.0632639,0.00014823866,0.9326429,0.003011555,0.000020156313],"about_ca_topic_score_codex":0.0064282916,"about_ca_topic_score_gemma":0.003429423,"teacher_disagreement_score":0.052185364,"about_ca_system_score_codex":0.0035563917,"about_ca_system_score_gemma":0.0039088055,"threshold_uncertainty_score":0.2759859},"labels":[],"label_agreement":null},{"id":"W4312579772","doi":"10.28995/2075-7182-2022-21-497-511","title":"Findings of the The RuATD Shared Task 2022 on Artificial Text Detection in Russian","year":2022,"lang":"en","type":"article","venue":"Computational Linguistics and Intellectual Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Air Canada","funders":"National Research University Higher School of Economics","keywords":"Computer science; Task (project management); Automatic summarization; Artificial intelligence; Natural language processing; Machine translation; Paraphrase; Text generation; Class (philosophy); Margin (machine learning); Binary classification; Machine learning; Support vector machine","score_opus":0.02515443092615935,"score_gpt":0.2439761048153211,"score_spread":0.21882167388916174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312579772","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6033532,0.019477528,0.06301678,0.01121063,0.007230356,0.002524397,0.20080094,0.037279155,0.05510703],"genre_scores_gemma":[0.46458483,0.001245266,0.049257442,0.00289898,0.0012911588,0.0019038966,0.45045477,0.003700035,0.0246636],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.96895784,0.016878111,0.001805745,0.0052465973,0.0052197673,0.0018919705],"domain_scores_gemma":[0.96412283,0.0160538,0.0011475267,0.007650309,0.007816317,0.0032092843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019936463,0.003834904,0.0030184372,0.0032206343,0.0033800513,0.0044149146,0.0032490382,0.0050078267,0.0064859795],"category_scores_gemma":[0.040242188,0.0008339607,0.0026613101,0.00216307,0.002020122,0.005294439,0.0090748435,0.0035465774,0.010616379],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0053831474,0.0047989236,0.024298701,0.005774525,0.0014858108,0.0015775958,0.005633694,0.019971058,0.022294138,0.0043334328,0.622107,0.28234193],"study_design_scores_gemma":[0.0031016443,0.005974512,0.16899887,0.0015344627,0.0019464407,0.00413678,0.009669654,0.17164537,0.07633659,0.019771945,0.5358699,0.0010138138],"about_ca_topic_score_codex":0.021250129,"about_ca_topic_score_gemma":0.027113449,"teacher_disagreement_score":0.021250129,"about_ca_system_score_codex":0.0026705526,"about_ca_system_score_gemma":0.0030148705,"threshold_uncertainty_score":0.10543531},"labels":[],"label_agreement":null},{"id":"W4312691443","doi":"10.14778/3554821.3554869","title":"SmartBench","year":2022,"lang":"en","type":"article","venue":"Proceedings of the VLDB Endowment","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Benchmark (surveying); Computer science; Question answering; Usability; Cover (algebra); Task (project management); Software deployment; Natural language; Information retrieval; Quality (philosophy); Natural language processing; Artificial intelligence; Software engineering; Human–computer interaction; Systems engineering; Engineering","score_opus":0.014372961389645013,"score_gpt":0.20559853763416233,"score_spread":0.19122557624451733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312691443","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012956902,0.0024833588,0.21339303,0.0013929046,0.000843357,0.001489803,0.122251615,0.57015604,0.07503302],"genre_scores_gemma":[0.07148253,0.0017575945,0.24619932,0.0017695762,0.00024308586,0.002089349,0.57529485,0.058703825,0.04245991],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971499,0.0007481547,0.00036457076,0.0007594156,0.0007755948,0.00020239853],"domain_scores_gemma":[0.993634,0.0029544844,0.00024699216,0.00155558,0.0013067517,0.0003022793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00327499,0.0018173169,0.0010354827,0.003552251,0.00068371807,0.0034657368,0.0045073424,0.0015957503,0.05794492],"category_scores_gemma":[0.014447127,0.0010860441,0.0015619916,0.003222921,0.00063701026,0.005446407,0.0038457667,0.0018066921,0.041099686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012416388,0.0003884413,0.0024806727,0.0026706385,0.00012484692,0.00046651333,0.0008909164,0.0056400816,0.008594096,0.03274715,0.69026244,0.25449258],"study_design_scores_gemma":[0.00038500482,0.00033154106,0.0019208969,0.00030153518,0.00007524446,0.00044279077,0.00038275294,0.0491253,0.01721393,0.030996472,0.8987173,0.000107189],"about_ca_topic_score_codex":0.004169803,"about_ca_topic_score_gemma":0.0052753477,"teacher_disagreement_score":0.05794492,"about_ca_system_score_codex":0.0011460982,"about_ca_system_score_gemma":0.001470749,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4312730985","doi":"10.1109/icpr56361.2022.9956503","title":"Embedded Spherical Topic Models for Supervised Learning","year":2022,"lang":"en","type":"article","venue":"2022 26th International Conference on Pattern Recognition (ICPR)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Discriminative model; Inference; Topic model; Artificial intelligence; Metadata; Probabilistic logic; Machine learning; Graph; Graphical model; Supervised learning; Information retrieval; Theoretical computer science; Artificial neural network","score_opus":0.12099668486804534,"score_gpt":0.3002750955334217,"score_spread":0.17927841066537636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312730985","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041702073,0.00040921717,0.9943387,0.00015778285,0.000023489842,0.000034793626,0.0001695068,0.00036045437,0.00033579988],"genre_scores_gemma":[0.48081714,0.0023552666,0.5016296,0.00057450903,0.00069431384,0.0011933008,0.0046573146,0.0006910917,0.007387529],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973822,0.0013219462,0.00013933913,0.0006578219,0.000361227,0.00013743475],"domain_scores_gemma":[0.99206066,0.0056049335,0.00063169864,0.00083941605,0.00071089144,0.00015242938],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004259909,0.0016165891,0.0019430306,0.0020684032,0.0006359151,0.0015003299,0.0029717064,0.0017125738,0.0024454377],"category_scores_gemma":[0.013477326,0.0009643009,0.0021779465,0.0024418833,0.0014137654,0.0033278766,0.002094871,0.0032858083,0.0014853496],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016588518,0.00012033684,0.0019429079,0.00024258038,0.00024456508,0.000109509645,0.00037350942,0.76018363,0.0014957292,0.10981626,0.0050056884,0.12029939],"study_design_scores_gemma":[0.000005666742,0.000010982437,0.000096776654,0.0000079868,0.0000073447877,0.000010088109,0.000010100634,0.96246886,0.00015346623,0.036647603,0.0005739424,0.0000071956274],"about_ca_topic_score_codex":0.006446209,"about_ca_topic_score_gemma":0.008592916,"teacher_disagreement_score":0.006446209,"about_ca_system_score_codex":0.0016651573,"about_ca_system_score_gemma":0.0013461027,"threshold_uncertainty_score":0.022528827},"labels":[],"label_agreement":null},{"id":"W4312839873","doi":"10.1007/978-981-19-8746-5_11","title":"Hierarchical Topic Model Inference by Community Discovery on Word Co-occurrence Networks","year":2022,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; University of Alberta","funders":"","keywords":"Topic model; Latent Dirichlet allocation; Computer science; Hierarchy; Inference; Set (abstract data type); Probabilistic logic; Graph; Community structure; Graphical model; Word (group theory); Artificial intelligence; Data science; Information retrieval; Natural language processing; Theoretical computer science; Mathematics","score_opus":0.071469626240898,"score_gpt":0.31919113455464243,"score_spread":0.24772150831374443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312839873","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009368832,0.0008679478,0.98776025,0.0002756732,0.000063707346,0.00004453661,0.00035576048,0.00054654194,0.0007168609],"genre_scores_gemma":[0.31559345,0.0021526539,0.66613024,0.00034634888,0.0007300629,0.0005359961,0.004977405,0.00063745514,0.008896401],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972493,0.001325576,0.00013573882,0.0007171738,0.0003936234,0.00017859641],"domain_scores_gemma":[0.9857691,0.012148199,0.00042731204,0.0009070081,0.0005593185,0.00018906403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035701492,0.0011346358,0.0020931899,0.003687625,0.0011715909,0.0022787752,0.003244674,0.0020430642,0.003110009],"category_scores_gemma":[0.01800203,0.0014240756,0.0023915193,0.005473704,0.001051955,0.004807973,0.002731976,0.003330106,0.0020826606],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064804516,0.00039403772,0.0071257213,0.0006761721,0.00088193244,0.0004831089,0.0009056003,0.3699554,0.00859667,0.09749811,0.023849865,0.48898533],"study_design_scores_gemma":[0.000011233742,0.000010537637,0.00035516382,0.000014480798,0.000033981378,0.000048933045,0.00002268454,0.9571869,0.00048049676,0.040706612,0.0011166903,0.000012349107],"about_ca_topic_score_codex":0.007055743,"about_ca_topic_score_gemma":0.00926719,"teacher_disagreement_score":0.007055743,"about_ca_system_score_codex":0.0011587241,"about_ca_system_score_gemma":0.0010305915,"threshold_uncertainty_score":0.018881023},"labels":[],"label_agreement":null},{"id":"W4312925950","doi":"10.2196/39077","title":"German Medical Named Entity Recognition Model and Data Set Creation Using Machine Translation and Word Alignment: Algorithm Development and Validation","year":2022,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Bildung und Forschung","keywords":"Computer science; Test set; Artificial intelligence; Machine translation; Natural language processing; Named-entity recognition; Set (abstract data type); Data set; Test data; German; Annotation; Data mining; Information retrieval; Programming language","score_opus":0.2205561870171737,"score_gpt":0.4362133696296889,"score_spread":0.21565718261251518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312925950","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2648811,0.0016797973,0.6872306,0.0017689002,0.0004924258,0.0017338713,0.0146560855,0.021773113,0.005784095],"genre_scores_gemma":[0.4451727,0.00063683785,0.4958836,0.00044894547,0.00008267788,0.0018267101,0.051567934,0.0005485035,0.0038321523],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978461,0.0007472348,0.0002486711,0.00065284496,0.00037483638,0.00013026384],"domain_scores_gemma":[0.9965328,0.0016344731,0.00021814257,0.0006570451,0.0008503898,0.00010722581],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004840751,0.0014181533,0.0006957667,0.0018548047,0.0007788754,0.0012792605,0.002173429,0.0016107345,0.0034017896],"category_scores_gemma":[0.010468024,0.00044924626,0.0014072158,0.0014946673,0.0006378094,0.0018104436,0.0018396422,0.0019741764,0.002335816],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092653086,0.00087890873,0.0135475,0.0008706579,0.0005722202,0.0010492159,0.00035155375,0.5489358,0.008765871,0.007964456,0.032655545,0.38348168],"study_design_scores_gemma":[0.00007510697,0.00013507708,0.0023425415,0.000056121193,0.00004134409,0.00018316461,0.00010226557,0.9780573,0.010713946,0.00316023,0.0050995285,0.00003340221],"about_ca_topic_score_codex":0.015666796,"about_ca_topic_score_gemma":0.014578762,"teacher_disagreement_score":0.015666796,"about_ca_system_score_codex":0.0017833285,"about_ca_system_score_gemma":0.0020937498,"threshold_uncertainty_score":0.031151175},"labels":[],"label_agreement":null},{"id":"W4313247739","doi":"10.23962/ajic.i30.13906","title":"A word embedding trained on South African news data","year":2022,"lang":"en","type":"article","venue":"The African Journal of Information and Communication (AJIC)","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Xanadu Quantum Technologies (Canada)","funders":"National Research Foundation; Department of Science and Innovation, South Africa","keywords":"Word embedding; Word2vec; Embedding; Word (group theory); Vocabulary; Computer science; Natural language processing; Artificial intelligence; Representation (politics); Information retrieval; Mathematics; Linguistics; Political science","score_opus":0.05390517709744123,"score_gpt":0.27291527227909523,"score_spread":0.219010095181654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313247739","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89507645,0.0030175382,0.074998945,0.0009989032,0.0010578877,0.0004768839,0.01229143,0.00395378,0.008128207],"genre_scores_gemma":[0.8599575,0.0011104171,0.08807345,0.00022733175,0.00019583153,0.0003126408,0.041108202,0.0002419913,0.008772618],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9992693,0.0003029258,0.00005772812,0.00018484311,0.000109810004,0.0000752829],"domain_scores_gemma":[0.99781847,0.0011513442,0.00008778314,0.0003648112,0.00051218324,0.00006538219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016277024,0.0010925198,0.00040549418,0.001058154,0.00034144655,0.00060455693,0.00039629327,0.0005536516,0.0018149914],"category_scores_gemma":[0.0061470554,0.00021477774,0.00066303636,0.0009158873,0.0003713695,0.0016546849,0.000840157,0.0012358286,0.0018515456],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022596815,0.0013829204,0.023997456,0.0012616003,0.0006077305,0.0006750356,0.0012352206,0.14700934,0.03975912,0.0033944747,0.049589433,0.7288279],"study_design_scores_gemma":[0.00018328309,0.0012771322,0.023010604,0.00020948681,0.00022897476,0.000511263,0.0011397841,0.8801639,0.054446675,0.003755838,0.034966893,0.00010613298],"about_ca_topic_score_codex":0.0052571604,"about_ca_topic_score_gemma":0.006842682,"teacher_disagreement_score":0.0052571604,"about_ca_system_score_codex":0.00045650732,"about_ca_system_score_gemma":0.00060942664,"threshold_uncertainty_score":0.010453105},"labels":[],"label_agreement":null},{"id":"W4313331018","doi":"10.14569/ijacsa.2022.0131234","title":"Transfer Learning for Closed Domain Question Answering in COVID-19","year":2022,"lang":"en","type":"article","venue":"International Journal of Advanced Computer Science and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Direktorat Riset and Pengembangan, Universitas Indonesia; Universitas Indonesia","keywords":"Computer science; Benchmark (surveying); Labrador Retriever; Cosine similarity; Transfer of learning; Baseline (sea); Question answering; Artificial intelligence; Coronavirus disease 2019 (COVID-19); Domain (mathematical analysis); Similarity (geometry); Open domain; Machine learning; Pattern recognition (psychology); Mathematics","score_opus":0.016209179709699137,"score_gpt":0.31620118342794273,"score_spread":0.2999920037182436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313331018","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.119497724,0.0034507527,0.80337423,0.0016833646,0.00049492245,0.0012726446,0.004607924,0.05638115,0.009237274],"genre_scores_gemma":[0.5003274,0.0006845094,0.46497062,0.0013993012,0.0002654684,0.000945201,0.02353413,0.0006453175,0.0072279572],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954189,0.0015523417,0.0003832789,0.0013130377,0.00095228094,0.00038019722],"domain_scores_gemma":[0.9937202,0.0032319955,0.00024830908,0.0011637919,0.0013490923,0.00028665367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006227191,0.0014172745,0.0017942582,0.0025734517,0.0009920913,0.001957005,0.0032366524,0.0024854797,0.006621757],"category_scores_gemma":[0.017455729,0.000528358,0.0013644587,0.0016780867,0.00090118224,0.0058356794,0.0044368035,0.0033208444,0.004507336],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001192982,0.0018788333,0.00610941,0.0013272539,0.00021710405,0.00072477217,0.0013369139,0.066496655,0.03339394,0.010988887,0.054936104,0.8213972],"study_design_scores_gemma":[0.00015581823,0.00046193713,0.0026043453,0.00005291023,0.000056534405,0.00032116342,0.00043177546,0.92566997,0.027542664,0.018640969,0.02397913,0.000082802035],"about_ca_topic_score_codex":0.007484756,"about_ca_topic_score_gemma":0.0056491974,"teacher_disagreement_score":0.007484756,"about_ca_system_score_codex":0.002181781,"about_ca_system_score_gemma":0.0020056178,"threshold_uncertainty_score":0.032932937},"labels":[],"label_agreement":null},{"id":"W4313343470","doi":"10.1007/978-3-662-66544-2_3","title":"Named Entity Recognition on CORD-19 Bio-Medical Dataset with Tolerance Rough Sets","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg; University of Alberta","funders":"","keywords":"Computer science; Named-entity recognition; Task (project management); Artificial intelligence; Natural language processing; Process (computing)","score_opus":0.033834887661863655,"score_gpt":0.27204638316298885,"score_spread":0.2382114955011252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313343470","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11458951,0.008200837,0.02395429,0.0017920663,0.0015181326,0.0008685591,0.82770133,0.0137892645,0.0075859986],"genre_scores_gemma":[0.058928676,0.0010166533,0.030166296,0.00024649114,0.0001353489,0.00038884615,0.90663993,0.000111720736,0.0023660497],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983026,0.00026823947,0.00026621166,0.00054624706,0.00040915876,0.00020753655],"domain_scores_gemma":[0.9983444,0.00042714944,0.00011927485,0.0004973735,0.00045419807,0.00015770263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014686253,0.0018064196,0.0016544675,0.0052501434,0.0009969391,0.0015289948,0.0018806771,0.0018580236,0.0042057913],"category_scores_gemma":[0.0037500104,0.00028465173,0.0020539733,0.004088035,0.00045888405,0.0010328213,0.0013476684,0.001132795,0.0050971215],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002064234,0.0009896505,0.031125676,0.0029766022,0.00092170166,0.0021269529,0.00020858127,0.01488185,0.009422936,0.0019833834,0.62392175,0.30937666],"study_design_scores_gemma":[0.0012209653,0.0017442931,0.123782575,0.0014382852,0.0020566336,0.0074224453,0.0020896334,0.23675579,0.044747025,0.0123071,0.5658322,0.0006031083],"about_ca_topic_score_codex":0.020081753,"about_ca_topic_score_gemma":0.030213,"teacher_disagreement_score":0.020081753,"about_ca_system_score_codex":0.0013543318,"about_ca_system_score_gemma":0.002712348,"threshold_uncertainty_score":0.039929748},"labels":[],"label_agreement":null},{"id":"W4313459318","doi":"10.1162/tacl_a_00529","title":"<scp>FaithDial</scp>: A Faithful Benchmark for Information-Seeking Dialogue","year":2022,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute; University of Alberta","funders":"Alberta Machine Intelligence Institute; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Benchmark (surveying); Computer science; Utterance; Hallucinating; Natural language processing; Artificial intelligence; Crowdsourcing; Machine learning; World Wide Web","score_opus":0.012802059413064245,"score_gpt":0.23327537375003826,"score_spread":0.220473314336974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313459318","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63366073,0.00894468,0.16756919,0.0037707202,0.0031100567,0.0025518616,0.055120267,0.086245805,0.039026666],"genre_scores_gemma":[0.75801027,0.00062405487,0.11138412,0.0010192979,0.00041177648,0.0011389221,0.11325579,0.002625508,0.011530277],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99071336,0.005456122,0.0004929999,0.0016526806,0.0012531548,0.00043166088],"domain_scores_gemma":[0.9827779,0.009158574,0.00065969565,0.0034600128,0.002586874,0.0013568354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069312924,0.0025942014,0.0010373964,0.002702551,0.0017454464,0.0027569104,0.0028962754,0.00326904,0.0062167747],"category_scores_gemma":[0.026048586,0.00046189944,0.0012759641,0.0012932185,0.0016966413,0.0031484384,0.004144596,0.0029489922,0.0044606905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006063689,0.004719162,0.017323036,0.005191254,0.0009193003,0.0015230889,0.0039719245,0.12583335,0.043824118,0.008591704,0.24759313,0.5344463],"study_design_scores_gemma":[0.0009337985,0.0039047725,0.020055862,0.0004995764,0.00020500069,0.0013770428,0.0032031864,0.7813051,0.07424955,0.016947102,0.09688436,0.00043474374],"about_ca_topic_score_codex":0.0092393,"about_ca_topic_score_gemma":0.014117369,"teacher_disagreement_score":0.0092393,"about_ca_system_score_codex":0.0017298006,"about_ca_system_score_gemma":0.0013571933,"threshold_uncertainty_score":0.036656618},"labels":[],"label_agreement":null},{"id":"W4313549837","doi":"10.1145/3551349.3556900","title":"AST-Probe: Recovering abstract syntax trees from hidden representations of pre-trained language models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; ENCODE; Natural language processing; Syntax; Artificial intelligence; Language model; Subspace topology; Representation (politics); Abstract syntax tree; Natural language","score_opus":0.028025031780247716,"score_gpt":0.27266450808684883,"score_spread":0.24463947630660113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313549837","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06485791,0.0002909461,0.91919637,0.0002500639,0.000081002814,0.00010174273,0.0015913883,0.01264947,0.0009811404],"genre_scores_gemma":[0.534902,0.00034719388,0.4455696,0.00027232128,0.000053227406,0.0005003145,0.013041134,0.0015493487,0.0037648708],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99938226,0.00018851462,0.000034727695,0.00018063986,0.00011626649,0.00009757549],"domain_scores_gemma":[0.9975764,0.001390454,0.0001523333,0.00047525886,0.00030998825,0.00009554937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009387343,0.0026364543,0.0008406129,0.0010771978,0.00045119063,0.0010330884,0.0015402385,0.0013188638,0.0029120494],"category_scores_gemma":[0.006681967,0.0007159693,0.0018094039,0.0008495739,0.00076781795,0.0032274507,0.0019641072,0.0033841517,0.0019502215],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007732698,0.0002670149,0.0048140525,0.0005835808,0.0002561339,0.00046891117,0.000901664,0.31234077,0.06711205,0.012939248,0.016493173,0.5830502],"study_design_scores_gemma":[0.000022570232,0.00009290837,0.0006597228,0.00002148198,0.000026025316,0.000060448496,0.000106237836,0.97291136,0.012516403,0.011944824,0.0016111562,0.00002682885],"about_ca_topic_score_codex":0.0054198187,"about_ca_topic_score_gemma":0.00756921,"teacher_disagreement_score":0.0054198187,"about_ca_system_score_codex":0.00065071485,"about_ca_system_score_gemma":0.0018487014,"threshold_uncertainty_score":0.01077652},"labels":[],"label_agreement":null},{"id":"W4313558932","doi":"10.3390/ai4010004","title":"End-to-End Transformer-Based Models in Textual-Based NLP","year":2023,"lang":"en","type":"article","venue":"AI","topic":"Topic Modeling","field":"Computer Science","cited_by":144,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Transformer; Computer science; Architecture; ENCODE; Artificial intelligence; Natural language processing; Language model; Machine learning; Engineering; Voltage; Electrical engineering","score_opus":0.039672393935301024,"score_gpt":0.27630558045827475,"score_spread":0.23663318652297372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313558932","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007227306,0.0010717928,0.982633,0.0005457551,0.00007531534,0.00010791523,0.0006074184,0.004565914,0.0031656104],"genre_scores_gemma":[0.43608385,0.0040071397,0.54236746,0.0007324124,0.00016190199,0.0005547076,0.004457169,0.0010998787,0.010535491],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99906427,0.00034368582,0.00007772753,0.00025662602,0.0001813983,0.00007632322],"domain_scores_gemma":[0.9964484,0.0023862903,0.00012050644,0.0004896768,0.00046317518,0.000092023605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022234144,0.0008903819,0.00079307595,0.0009859641,0.0005342563,0.002275536,0.0022599385,0.0013113748,0.0055127395],"category_scores_gemma":[0.008334092,0.0005224588,0.0011787849,0.0014318165,0.00091744063,0.0057056583,0.0017879358,0.0021457921,0.0031082944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006082999,0.00023597565,0.0022486933,0.00095953245,0.0002189625,0.00045090396,0.0009820694,0.30627388,0.009115397,0.09715096,0.01690299,0.56485236],"study_design_scores_gemma":[0.000017829652,0.000046974474,0.00018109937,0.000051171723,0.000040283336,0.00010086209,0.000069458605,0.9398791,0.0032610747,0.049858067,0.0064761387,0.000017899692],"about_ca_topic_score_codex":0.008226775,"about_ca_topic_score_gemma":0.01028452,"teacher_disagreement_score":0.008226775,"about_ca_system_score_codex":0.0013504107,"about_ca_system_score_gemma":0.0015389199,"threshold_uncertainty_score":0.018441975},"labels":[],"label_agreement":null},{"id":"W4313563640","doi":"10.1145/3551349.3556917","title":"Automatic Comment Generation via Multi-Pass Deliberation","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China; Youth Innovation Promotion Association; Chinese Academy of Sciences; National Science Foundation","keywords":"Computer science; Deliberation; Code (set theory); Iterative and incremental development; Process (computing); Python (programming language); Java; Source lines of code; Programming language; Artificial intelligence; Software; Software engineering","score_opus":0.04805307435243965,"score_gpt":0.2612893608342793,"score_spread":0.21323628648183962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313563640","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023117915,0.00063928997,0.9173836,0.00086247374,0.00034979556,0.0012924712,0.0021429108,0.050528295,0.003683271],"genre_scores_gemma":[0.17627025,0.00039754855,0.7941882,0.00058323523,0.00031322622,0.0016776522,0.0117605375,0.0036150324,0.011194353],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98656785,0.0058502704,0.0009493489,0.0028433762,0.0032609901,0.0005281746],"domain_scores_gemma":[0.9573982,0.025596956,0.0022214607,0.0055663246,0.008093831,0.0011233257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008423876,0.002811983,0.0019105496,0.003447601,0.0013482776,0.0026930675,0.0028913484,0.0019721827,0.008005735],"category_scores_gemma":[0.044751663,0.0007102953,0.0022671,0.0018383056,0.0011877188,0.005187557,0.0052947504,0.0023783478,0.0084892195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001561663,0.0005101518,0.0071260636,0.0020178703,0.0002698627,0.00091641257,0.0057538846,0.0124588655,0.041455172,0.010373104,0.07360107,0.84395593],"study_design_scores_gemma":[0.0005701978,0.0007222801,0.0052489126,0.00046004012,0.00036123014,0.00096620474,0.003150908,0.706037,0.096203804,0.051239267,0.13456644,0.0004737529],"about_ca_topic_score_codex":0.0024001992,"about_ca_topic_score_gemma":0.0032292919,"teacher_disagreement_score":0.008423876,"about_ca_system_score_codex":0.00103159,"about_ca_system_score_gemma":0.002282677,"threshold_uncertainty_score":0.04455024},"labels":[],"label_agreement":null},{"id":"W4315641315","doi":"10.31219/osf.io/naf8t","title":"Reproducible research practices and transparency across linguistics","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"","keywords":"Transparency (behavior); Applied linguistics; Open science; Publishing; Data sharing; Empirical research; Best practice; Data science; Psychology; Computer science; Engineering ethics; Sociology; Political science; Linguistics; Engineering; Medicine; Statistics; Mathematics","score_opus":0.5162929941209532,"score_gpt":0.5143864763262007,"score_spread":0.0019065177947524825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4315641315","genre_codex":"methods","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14119603,0.021824026,0.5720747,0.19336528,0.004530591,0.0113203395,0.0022973532,0.002170546,0.05122119],"genre_scores_gemma":[0.70689046,0.0036321327,0.25925875,0.011908669,0.0017706784,0.011194938,0.0011065452,0.00081225927,0.003425648],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.12087328,0.69306207,0.073909715,0.032353822,0.0747927,0.0050083846],"domain_scores_gemma":[0.03157609,0.59361756,0.08205645,0.22532815,0.063288495,0.0041332874],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.72947866,0.001795186,0.004227906,0.01766392,0.010990032,0.029075272,0.009062061,0.007960334,0.005235071],"category_scores_gemma":[0.86789215,0.0034449855,0.0034060876,0.016053563,0.039221708,0.034741044,0.022947302,0.011264979,0.0022171363],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010610335,0.0007362833,0.11336726,0.014186451,0.0018242915,0.0015717805,0.25203562,0.002711076,0.0069911997,0.23514079,0.021348622,0.34902558],"study_design_scores_gemma":[0.00073088374,0.0010774945,0.072808556,0.030558268,0.0008324461,0.0019869478,0.061903045,0.007175587,0.012663152,0.60517675,0.20422955,0.0008573168],"about_ca_topic_score_codex":0.0054061026,"about_ca_topic_score_gemma":0.0044099544,"teacher_disagreement_score":0.27052134,"about_ca_system_score_codex":0.010888828,"about_ca_system_score_gemma":0.05017671,"threshold_uncertainty_score":0.33360106},"labels":[],"label_agreement":null},{"id":"W4315796871","doi":"10.31222/osf.io/rcews","title":"Reproducible research practices and transparency across linguistics","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"","keywords":"Transparency (behavior); Applied linguistics; Open science; Publishing; Empirical research; Data sharing; Best practice; Data science; Psychology; Sociology; Computer science; Public relations; Political science; Linguistics; Medicine; Statistics; Mathematics","score_opus":0.5162929941209532,"score_gpt":0.5143864763262007,"score_spread":0.0019065177947524825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4315796871","genre_codex":"methods","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14119603,0.021824026,0.5720747,0.19336528,0.004530591,0.0113203395,0.0022973532,0.002170546,0.05122119],"genre_scores_gemma":[0.70689046,0.0036321327,0.25925875,0.011908669,0.0017706784,0.011194938,0.0011065452,0.00081225927,0.003425648],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.12087328,0.69306207,0.073909715,0.032353822,0.0747927,0.0050083846],"domain_scores_gemma":[0.03157609,0.59361756,0.08205645,0.22532815,0.063288495,0.0041332874],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.72947866,0.001795186,0.004227906,0.01766392,0.010990032,0.029075272,0.009062061,0.007960334,0.005235071],"category_scores_gemma":[0.86789215,0.0034449855,0.0034060876,0.016053563,0.039221708,0.034741044,0.022947302,0.011264979,0.0022171363],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010610335,0.0007362833,0.11336726,0.014186451,0.0018242915,0.0015717805,0.25203562,0.002711076,0.0069911997,0.23514079,0.021348622,0.34902558],"study_design_scores_gemma":[0.00073088374,0.0010774945,0.072808556,0.030558268,0.0008324461,0.0019869478,0.061903045,0.007175587,0.012663152,0.60517675,0.20422955,0.0008573168],"about_ca_topic_score_codex":0.0054061026,"about_ca_topic_score_gemma":0.0044099544,"teacher_disagreement_score":0.27052134,"about_ca_system_score_codex":0.010888828,"about_ca_system_score_gemma":0.05017671,"threshold_uncertainty_score":0.33360106},"labels":[],"label_agreement":null},{"id":"W4316041138","doi":"10.20944/preprints202301.0219.v1","title":"When to Use Large Language Model: Upper Bound Analysis of BM25 Algorithms in Reading Comprehension Task","year":2023,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Geomechanica (Canada)","funders":"","keywords":"Computer science; Task (project management); Representation (politics); Reading (process); Comprehension; Natural language processing; Artificial intelligence; Language model; Language understanding; Machine learning; Upper and lower bounds; Linguistics; Engineering; Mathematics; Programming language","score_opus":0.15436600947938695,"score_gpt":0.36636229905745393,"score_spread":0.211996289578067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4316041138","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43344247,0.02865491,0.48594174,0.012368135,0.0011099684,0.00059751555,0.0032207966,0.007139536,0.027525043],"genre_scores_gemma":[0.843441,0.0014504822,0.1419375,0.0014620835,0.00054549077,0.00042426746,0.0048417114,0.0010157896,0.0048816474],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99285394,0.0038422171,0.00037192003,0.0013385004,0.0010598649,0.0005335766],"domain_scores_gemma":[0.95764095,0.035461802,0.0009464374,0.0029162169,0.0021433532,0.0008912809],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017106496,0.0016688998,0.0019004632,0.0026494006,0.0014696053,0.0044281175,0.002511418,0.0036243899,0.0036593005],"category_scores_gemma":[0.0697031,0.00060694484,0.0010883918,0.002045172,0.0010932544,0.0076460806,0.0019442998,0.004187285,0.0019094334],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002744706,0.0011477112,0.019701123,0.0009397688,0.0007634119,0.00027960166,0.00071432977,0.16864401,0.0073481086,0.019301571,0.067651674,0.71076405],"study_design_scores_gemma":[0.00011475593,0.0002878794,0.008542614,0.00011612082,0.00015283939,0.00010869547,0.0003078071,0.9560197,0.004136437,0.02699953,0.0031315137,0.00008203677],"about_ca_topic_score_codex":0.011562705,"about_ca_topic_score_gemma":0.015304766,"teacher_disagreement_score":0.017106496,"about_ca_system_score_codex":0.002480426,"about_ca_system_score_gemma":0.002243298,"threshold_uncertainty_score":0.090468824},"labels":[],"label_agreement":null},{"id":"W4316135769","doi":"10.48550/arxiv.2301.05220","title":"Adversarial Adaptation for French Named Entity Recognition","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Mitacs","keywords":"Overfitting; Computer science; Named-entity recognition; Domain adaptation; Transformer; Adversarial system; Artificial intelligence; Natural language processing; Machine learning; Adaptation (eye); Task (project management); Labeled data; Artificial neural network; Classifier (UML)","score_opus":0.22014513095720462,"score_gpt":0.2070984495315644,"score_spread":0.013046681425640222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4316135769","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038525857,0.00052760245,0.9546803,0.00048597116,0.00009428436,0.000059401602,0.00038227794,0.0018846067,0.0033596107],"genre_scores_gemma":[0.871887,0.00048614177,0.114739016,0.0005061408,0.00012771787,0.00019825372,0.0021395583,0.00031357753,0.009602663],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991948,0.0003765071,0.000026621858,0.00023024878,0.00010234655,0.000069415255],"domain_scores_gemma":[0.9979321,0.0014197329,0.000116759766,0.00031294787,0.00017493118,0.00004372837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018556038,0.00093230413,0.000701324,0.0005597411,0.00035843888,0.00057996786,0.0010357843,0.0008103538,0.0017391831],"category_scores_gemma":[0.004621748,0.00031961408,0.0006892646,0.00067191594,0.0008778919,0.0010496805,0.0011743378,0.0015466978,0.0008777505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000115260984,0.00004340156,0.0008603124,0.000036544294,0.000058588557,0.00012832934,0.00008710709,0.9194502,0.0027479227,0.010769891,0.004919369,0.06078309],"study_design_scores_gemma":[0.0000023026528,0.0000063583557,0.00012158407,0.0000023580408,0.0000031837242,0.000013105045,0.0000046546324,0.99554247,0.0005527139,0.0032486287,0.0004983803,0.000004276687],"about_ca_topic_score_codex":0.0060200994,"about_ca_topic_score_gemma":0.005423038,"teacher_disagreement_score":0.0060200994,"about_ca_system_score_codex":0.00086659513,"about_ca_system_score_gemma":0.000417236,"threshold_uncertainty_score":0.011970103},"labels":[],"label_agreement":null},{"id":"W4317467432","doi":"10.1101/2023.01.18.524571","title":"Ensemble of deep learning language models to support the creation of living systematic reviews for the COVID-19 literature","year":2023,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institutes of Health Research; Horizon 2020 Framework Programme; Innosuisse - Schweizerische Agentur für Innovationsförderung; European Commission; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; National Science Foundation","keywords":"Artificial intelligence; Machine learning; Computer science; Ensemble forecasting; Ranking (information retrieval); Task (project management); Natural language processing; Class (philosophy); Ensemble learning; Systematic review; Data science; Information retrieval; MEDLINE","score_opus":0.053016200343038826,"score_gpt":0.2845664423135185,"score_spread":0.2315502419704797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4317467432","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24230772,0.080527455,0.60026556,0.014123486,0.00209448,0.0040721004,0.027864765,0.022223314,0.006521135],"genre_scores_gemma":[0.5264735,0.0068752887,0.43673602,0.0028248949,0.0005488677,0.0025121637,0.021165477,0.0004347944,0.0024289887],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9889163,0.0067426,0.001594855,0.001566671,0.0009741853,0.00020534774],"domain_scores_gemma":[0.94855833,0.038718365,0.0038411971,0.0026148665,0.0055455137,0.00072172994],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026405187,0.0018779795,0.002007566,0.006623822,0.0006831442,0.0027333016,0.0022508833,0.0016599274,0.0015542768],"category_scores_gemma":[0.069338724,0.00079939497,0.003757591,0.0028102908,0.00038959866,0.0027721054,0.0020343938,0.001985269,0.0012880592],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019197032,0.00067555386,0.04978055,0.013290242,0.008876809,0.00094832614,0.0010688655,0.15732853,0.006419004,0.0028426312,0.039582126,0.71726763],"study_design_scores_gemma":[0.00043032932,0.0006579785,0.007142331,0.0029256113,0.0036885557,0.00038142118,0.00025620134,0.95313954,0.0042800787,0.012092647,0.014871955,0.00013329112],"about_ca_topic_score_codex":0.008044721,"about_ca_topic_score_gemma":0.027621588,"teacher_disagreement_score":0.9735948,"about_ca_system_score_codex":0.0019985284,"about_ca_system_score_gemma":0.006215278,"threshold_uncertainty_score":0.13964564},"labels":[],"label_agreement":null},{"id":"W4317897852","doi":"10.1162/tacl_a_00539","title":"Cross-Lingual Dialogue Dataset Creation via Outline-Based Generation","year":2023,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Computer science; Naturalness; Natural language processing; Machine translation; Artificial intelligence; Modular design; Process (computing); Annotation; Benchmark (surveying); Information retrieval; Programming language","score_opus":0.039882268698897924,"score_gpt":0.31943519721844155,"score_spread":0.2795529285195436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4317897852","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1761277,0.0027197285,0.49184614,0.0019973363,0.0023476917,0.0037198788,0.19048384,0.09455238,0.03620532],"genre_scores_gemma":[0.2140879,0.00029963805,0.36212668,0.00057752256,0.00016969808,0.0037124245,0.40910858,0.0029006333,0.0070169396],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99248356,0.0030426825,0.0007109788,0.0023345817,0.0011060253,0.00032221264],"domain_scores_gemma":[0.9879822,0.0036359418,0.00057490537,0.003953732,0.0031680206,0.00068521785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057684584,0.0015601098,0.0008414801,0.0030258563,0.00195977,0.002244485,0.0026670469,0.0015489535,0.005085057],"category_scores_gemma":[0.018401427,0.0004998351,0.0016040136,0.001702635,0.001114367,0.0028563521,0.006356001,0.002460239,0.0051294663],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002467707,0.0017418885,0.020375872,0.004620685,0.00056692126,0.0018570736,0.010099903,0.029430669,0.06755533,0.019004676,0.31922954,0.5230498],"study_design_scores_gemma":[0.0005220053,0.0007840219,0.032792415,0.00087971834,0.00028119577,0.0014665755,0.0077177044,0.22366853,0.07518837,0.020573758,0.6355611,0.0005645612],"about_ca_topic_score_codex":0.0057424577,"about_ca_topic_score_gemma":0.011217095,"teacher_disagreement_score":0.0057684584,"about_ca_system_score_codex":0.001149979,"about_ca_system_score_gemma":0.0020503406,"threshold_uncertainty_score":0.030506909},"labels":[],"label_agreement":null},{"id":"W4318195334","doi":"10.1186/s12911-023-02117-3","title":"Entity and relation extraction from clinical case reports of COVID-19: a natural language processing approach","year":2023,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Public Health Ontario","funders":"Institute of Health Services and Policy Research; Canadian Institutes of Health Research","keywords":"Coronavirus disease 2019 (COVID-19); Health informatics; Computer science; Relation (database); Natural language processing; 2019-20 coronavirus outbreak; Information extraction; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Relationship extraction; Natural language; Artificial intelligence; Medicine; Data mining; Pathology; Public health","score_opus":0.08397814884412258,"score_gpt":0.4160037932922904,"score_spread":0.33202564444816784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318195334","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06856008,0.005497994,0.83149964,0.0051618074,0.00043590186,0.0029695127,0.062991634,0.017203435,0.005680041],"genre_scores_gemma":[0.13705458,0.0021036556,0.7854437,0.00045796303,0.00032345182,0.0010203199,0.07173539,0.0003160807,0.0015448913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948761,0.0013387116,0.0013370772,0.0012260048,0.0010549508,0.000167254],"domain_scores_gemma":[0.97502667,0.017189367,0.002959427,0.0021356689,0.0022560363,0.00043281642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006026203,0.0014795836,0.0010159715,0.016093155,0.0009552817,0.003721942,0.0018882783,0.0017028222,0.0032669525],"category_scores_gemma":[0.022111021,0.00063483114,0.002397446,0.0071958997,0.0008252191,0.0036061294,0.0026014464,0.0016802874,0.0025849396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005692678,0.0007354406,0.04357887,0.0064223334,0.0007570285,0.009608016,0.0030478819,0.015366957,0.04971684,0.012442315,0.05368128,0.8040737],"study_design_scores_gemma":[0.00031604563,0.000614876,0.08758153,0.0028035026,0.001970512,0.017931726,0.007362391,0.41669542,0.111442834,0.081625685,0.27112827,0.0005272478],"about_ca_topic_score_codex":0.0038444991,"about_ca_topic_score_gemma":0.00561336,"teacher_disagreement_score":0.016093155,"about_ca_system_score_codex":0.0011239155,"about_ca_system_score_gemma":0.0032433555,"threshold_uncertainty_score":0.031870008},"labels":[],"label_agreement":null},{"id":"W4318239436","doi":"10.48550/arxiv.2301.10493","title":"From Baseline to Top Performer: A Reproducibility Study of Approaches at the TREC 2021 Conversational Assistance Track","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Norges Forskningsråd; University of Waterloo","keywords":"Computer science; Baseline (sea); Pipeline (software); Information retrieval; Set (abstract data type); Margin (machine learning); Measure (data warehouse); Track (disk drive); Data mining; Machine learning","score_opus":0.2084443694460633,"score_gpt":0.22112230941693917,"score_spread":0.012677939970875879,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318239436","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79387325,0.01751921,0.11776646,0.0049925954,0.0049729245,0.0019599805,0.016995464,0.01662716,0.025292872],"genre_scores_gemma":[0.93547225,0.0007035139,0.036760394,0.0008337418,0.00088460033,0.00075459585,0.016743995,0.0031617817,0.004685131],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.8487178,0.08447782,0.011715784,0.024690963,0.027290408,0.0031072274],"domain_scores_gemma":[0.6571738,0.1568911,0.018284053,0.096227184,0.0651237,0.0063001523],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.11605639,0.0021785402,0.0018876069,0.004526152,0.003642777,0.0061672437,0.003726289,0.0025905864,0.0019510234],"category_scores_gemma":[0.24433634,0.0009346714,0.0022571534,0.003886826,0.0025492883,0.0063892435,0.0053449366,0.003497346,0.0038437578],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010401073,0.0031283533,0.2547584,0.00628782,0.007729565,0.0012086559,0.0192718,0.017837971,0.047100607,0.0049996856,0.119007245,0.50826883],"study_design_scores_gemma":[0.001602024,0.01714965,0.5544614,0.0015084359,0.0048609185,0.0041932403,0.01876961,0.1257832,0.11535754,0.018667815,0.13582261,0.0018234826],"about_ca_topic_score_codex":0.006933659,"about_ca_topic_score_gemma":0.008013124,"teacher_disagreement_score":0.8839436,"about_ca_system_score_codex":0.0026606761,"about_ca_system_score_gemma":0.002591354,"threshold_uncertainty_score":0.61377215},"labels":[{"model":"gemma","categories":["metaresearch"],"domain":"reproducibility","study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["metaresearch"],"domain":"reproducibility","study_design":"design_other","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"medium"}],"label_agreement":"split"},{"id":"W4318351452","doi":"10.48550/arxiv.2301.11305","title":"Cognitive Constraint Simulation and the Geometry of Human Authorship: A First-Principles Theory of AI Text Detection","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":154,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Discriminative model; Zero (linguistics); Computer science; Artificial intelligence; Curvature; Shot (pellet); Classifier (UML); Function (biology); Machine learning; Natural language processing; Algorithm; Pattern recognition (psychology); Mathematics; Geometry; Linguistics","score_opus":0.15016658314371803,"score_gpt":0.24122364926403164,"score_spread":0.09105706612031361,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318351452","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045131292,0.0005560985,0.93361443,0.0035673794,0.000079431215,0.00008197273,0.00020474126,0.00035146638,0.016413176],"genre_scores_gemma":[0.8408199,0.0005982576,0.15216766,0.00042310913,0.00021313383,0.00031018362,0.00025213612,0.00019256065,0.005023003],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99829966,0.00081333145,0.00006039352,0.00036606303,0.00035264692,0.0001079903],"domain_scores_gemma":[0.9864108,0.0097277835,0.0012275066,0.0015120534,0.0006678941,0.000453993],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027130742,0.00059688237,0.0010919471,0.0022727642,0.0013220425,0.0037673977,0.0022456665,0.0020035114,0.0052082585],"category_scores_gemma":[0.024964673,0.0006892755,0.00145727,0.0012234463,0.006318411,0.0065113846,0.002466881,0.0019788833,0.0007182327],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048697908,0.00004048426,0.0029132396,0.00018024391,0.000077973265,0.00016484143,0.00086439593,0.14659142,0.0008968378,0.8219785,0.0020606883,0.02418264],"study_design_scores_gemma":[0.000012178097,0.000018993413,0.00061898347,0.000020432848,0.000006366846,0.00007089341,0.00006722643,0.37674773,0.00020419045,0.6210875,0.0011224399,0.00002291246],"about_ca_topic_score_codex":0.0062367655,"about_ca_topic_score_gemma":0.0031043065,"teacher_disagreement_score":0.0062367655,"about_ca_system_score_codex":0.0026355467,"about_ca_system_score_gemma":0.0015624311,"threshold_uncertainty_score":0.019122362},"labels":[],"label_agreement":null},{"id":"W4318477386","doi":"10.1145/3580508","title":"Heterogeneous Graph Transformer for Meta-structure Learning with Application in Text Classification","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on the Web","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"Natural Science Foundation of Hebei Province","keywords":"Computer science; Graph; Theoretical computer science; Meta learning (computer science); Artificial intelligence; Extractor; Heterogeneous network; Metamodeling; Knowledge graph; Machine learning; Wireless network; Programming language","score_opus":0.048388398918130994,"score_gpt":0.2623894346928376,"score_spread":0.2140010357747066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318477386","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012592843,0.0006751535,0.9825891,0.0002337231,0.0000523464,0.00008515829,0.00047461165,0.0021932146,0.0011038766],"genre_scores_gemma":[0.40159068,0.0011798934,0.58738405,0.00045147905,0.00014835656,0.0003787679,0.004056412,0.00054618734,0.0042641577],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999335,0.0001883587,0.000037305323,0.0002618558,0.00012970876,0.000047703586],"domain_scores_gemma":[0.99848026,0.00079949736,0.00016416922,0.00028424172,0.00019824224,0.00007356982],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009970289,0.0012350566,0.0007563633,0.003957542,0.0005900003,0.0008518598,0.0018505703,0.00125803,0.0028070654],"category_scores_gemma":[0.0045739207,0.0004741844,0.0015731009,0.0036552364,0.0008493796,0.00373836,0.0011570135,0.0019115272,0.0011277045],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028963416,0.00027048157,0.0037708343,0.0005075797,0.00023133199,0.0003700357,0.00038316756,0.24777156,0.014168474,0.04922816,0.011533152,0.67147565],"study_design_scores_gemma":[0.000014146936,0.00003903439,0.00040741952,0.000022435392,0.000040132538,0.000073756135,0.000041815292,0.94172317,0.002236258,0.05231292,0.003076963,0.000011939461],"about_ca_topic_score_codex":0.0041265227,"about_ca_topic_score_gemma":0.007529917,"teacher_disagreement_score":0.0041265227,"about_ca_system_score_codex":0.0012706176,"about_ca_system_score_gemma":0.00090093666,"threshold_uncertainty_score":0.009390533},"labels":[],"label_agreement":null},{"id":"W4318606113","doi":"10.1109/ssci51031.2022.10022116","title":"Improving Topic Quality with Interactive Beta-Liouville Mixture Allocation Model","year":2022,"lang":"en","type":"article","venue":"2022 IEEE Symposium Series on Computational Intelligence (SSCI)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Latent Dirichlet allocation; Topic model; Computer science; Inference; Cluster analysis; Categorization; Artificial intelligence; Dirichlet distribution; Natural language processing; Task (project management); Machine learning; Set (abstract data type); Quality (philosophy); Hierarchical Dirichlet process; BETA (programming language); Mathematics","score_opus":0.02637901145546341,"score_gpt":0.28056295123112657,"score_spread":0.25418393977566317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318606113","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017021736,0.00051232375,0.9796708,0.00023579651,0.00005202465,0.000063001484,0.00006614479,0.0014407128,0.0009373673],"genre_scores_gemma":[0.55356187,0.00086834014,0.43717772,0.00058159104,0.00026523342,0.00051017734,0.00110938,0.00090226025,0.0050233696],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99686414,0.0015427949,0.00018280486,0.0006116936,0.00055775704,0.00024081758],"domain_scores_gemma":[0.992708,0.0047623427,0.00035093998,0.00086232036,0.0010284046,0.00028797888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065073296,0.0012214088,0.002115044,0.001814728,0.001106233,0.002677516,0.0034104993,0.002435758,0.002357789],"category_scores_gemma":[0.017098343,0.00093795074,0.0019273389,0.001885043,0.001149956,0.0054606446,0.0029301124,0.003153789,0.0020250308],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013383983,0.00066749257,0.0078603355,0.0002760802,0.00036575715,0.00022633902,0.001432741,0.5077273,0.015444934,0.033206753,0.009338727,0.42211524],"study_design_scores_gemma":[0.000032204985,0.000029879026,0.00023868447,0.0000068058716,0.000021407774,0.00003276194,0.000027572221,0.9898545,0.0014222597,0.007397053,0.000919444,0.000017438608],"about_ca_topic_score_codex":0.006553198,"about_ca_topic_score_gemma":0.007148977,"teacher_disagreement_score":0.006553198,"about_ca_system_score_codex":0.001521844,"about_ca_system_score_gemma":0.0016045949,"threshold_uncertainty_score":0.03441441},"labels":[],"label_agreement":null},{"id":"W4318710503","doi":"10.1145/3582900.3582911","title":"Report on the 16th Round of NII Testbeds and Community for Information Access Research (NTCIR-16)","year":2022,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Automatic summarization; Scope (computer science); Task (project management); Question answering; World Wide Web; Information retrieval; Programming language","score_opus":0.15888139140801708,"score_gpt":0.38471710326605507,"score_spread":0.225835711858038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4318710503","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01993208,0.008864041,0.04799532,0.08604194,0.06715492,0.023588004,0.522354,0.02389812,0.20017159],"genre_scores_gemma":[0.031308003,0.0021890968,0.064043246,0.0120690325,0.0044360976,0.022502325,0.63699317,0.0070898915,0.21936914],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9782559,0.0070191035,0.000775185,0.0018116339,0.008000305,0.004137963],"domain_scores_gemma":[0.9098712,0.0061392994,0.001973355,0.0069952905,0.035535935,0.039484892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.051834594,0.0027042455,0.002320431,0.004807806,0.0056541874,0.009656011,0.004628505,0.004573567,0.10563992],"category_scores_gemma":[0.035693344,0.0012171546,0.0019991295,0.0036230397,0.0013963943,0.007655884,0.017291471,0.0058908258,0.11361781],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034787142,0.00019534245,0.0008081992,0.00022864826,0.000023977436,0.0000633727,0.00018635849,0.00017778178,0.0009980357,0.0014029444,0.98169965,0.013867769],"study_design_scores_gemma":[0.00031121893,0.0003696176,0.006088213,0.00023182438,0.00004214864,0.000058608068,0.000735931,0.00063266023,0.0017655143,0.0022345325,0.9874418,0.00008798306],"about_ca_topic_score_codex":0.03310411,"about_ca_topic_score_gemma":0.05075531,"teacher_disagreement_score":0.10563992,"about_ca_system_score_codex":0.005045187,"about_ca_system_score_gemma":0.022850163,"threshold_uncertainty_score":0.35340077},"labels":[],"label_agreement":null},{"id":"W4319062633","doi":"10.1093/database/baac108","title":"Automatic Extraction of Medication Mentions from Tweets—Overview of the BioCreative VII Shared Task 3 Competition","year":2023,"lang":"en","type":"article","venue":"Database","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"U.S. National Library of Medicine","keywords":"Timeline; Computer science; Task (project management); Natural language processing; Artificial intelligence; Class (philosophy); Information retrieval; World Wide Web","score_opus":0.06896966406581621,"score_gpt":0.32966677997371774,"score_spread":0.2606971159079015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319062633","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23910318,0.018506227,0.101561256,0.014266493,0.0072061825,0.009696382,0.49079096,0.08224292,0.036626376],"genre_scores_gemma":[0.04804764,0.0012855891,0.09225586,0.0017127909,0.0006902334,0.005232583,0.83587146,0.0048865033,0.010017293],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98589176,0.0043043764,0.0013714862,0.0035820305,0.003554755,0.0012956532],"domain_scores_gemma":[0.9817442,0.0070872274,0.000913395,0.002997856,0.005191868,0.0020653855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017024947,0.005010755,0.0035256306,0.0073593785,0.0054556644,0.005365973,0.0057120137,0.0048898514,0.0084301075],"category_scores_gemma":[0.022108763,0.0015204897,0.0037105952,0.0054175216,0.0017267723,0.004106556,0.00931365,0.0042882487,0.016842458],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027895952,0.0018156077,0.00889563,0.0049155112,0.00070593756,0.0017740236,0.0021528124,0.0038457206,0.048015453,0.0021121784,0.7245072,0.19847038],"study_design_scores_gemma":[0.002447852,0.0020241085,0.051277254,0.0008238924,0.0007657585,0.0024385334,0.0044012433,0.06941837,0.07171553,0.007688433,0.78636247,0.00063655357],"about_ca_topic_score_codex":0.019928603,"about_ca_topic_score_gemma":0.03728904,"teacher_disagreement_score":0.019928603,"about_ca_system_score_codex":0.003681621,"about_ca_system_score_gemma":0.0080420915,"threshold_uncertainty_score":0.090037584},"labels":[],"label_agreement":null},{"id":"W4319301677","doi":"10.1162/neco_a_01563","title":"Large Language Models and the Reverse Turing Test","year":2023,"lang":"en","type":"article","venue":"Neural Computation","topic":"Topic Modeling","field":"Computer Science","cited_by":164,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of Biomedical Imaging and Bioengineering; Canadian Institute for Advanced Research; National Science Foundation","keywords":"Psychology; Cognitive psychology; Computer science; Cognitive science","score_opus":0.030632925800254278,"score_gpt":0.27425707848986514,"score_spread":0.24362415268961085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319301677","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16340198,0.011993247,0.6068138,0.10243322,0.0025289801,0.0006571512,0.0017656675,0.004003155,0.10640286],"genre_scores_gemma":[0.82440394,0.0021554963,0.14667068,0.008717726,0.0017994121,0.0012906203,0.0018615372,0.0010057662,0.012094664],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9473379,0.040723246,0.001906686,0.0042243986,0.004491286,0.0013164241],"domain_scores_gemma":[0.7187572,0.24087343,0.0048400667,0.023718407,0.00903827,0.002772656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03664291,0.0013639942,0.0024845067,0.0017956432,0.0025052964,0.0063967253,0.0032721192,0.0037275532,0.0071318885],"category_scores_gemma":[0.19619034,0.0008982833,0.0025008982,0.0010468429,0.012821618,0.01532806,0.0077655427,0.009220462,0.0020438237],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039581052,0.00011877215,0.0046875817,0.00069868873,0.00030307355,0.00073742564,0.0024365073,0.019533027,0.0005974481,0.8816514,0.021166751,0.0676735],"study_design_scores_gemma":[0.000054043885,0.00004180288,0.0004858665,0.000097799726,0.000025274481,0.00016320682,0.00023290662,0.03448628,0.00053355633,0.9568704,0.00695741,0.000051427298],"about_ca_topic_score_codex":0.0024178305,"about_ca_topic_score_gemma":0.0014224409,"teacher_disagreement_score":0.03664291,"about_ca_system_score_codex":0.0030940394,"about_ca_system_score_gemma":0.0026154846,"threshold_uncertainty_score":0.19378853},"labels":[],"label_agreement":null},{"id":"W4319779582","doi":"10.1109/icdmw58026.2022.00064","title":"Feature Extraction and Prediction of Combined Text and Survey Data using Two-Staged Modeling","year":2022,"lang":"en","type":"article","venue":"2022 IEEE International Conference on Data Mining Workshops (ICDMW)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Computer science; Artificial intelligence; Classifier (UML); Domain adaptation; Convolutional neural network; Feature extraction; Deep learning; Machine learning; Recurrent neural network; Artificial neural network; Pattern recognition (psychology); Data mining","score_opus":0.2961297511315246,"score_gpt":0.37188543190658574,"score_spread":0.07575568077506112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319779582","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18360691,0.0006053453,0.8027623,0.00061229523,0.00015450628,0.00020208592,0.004978412,0.0052582053,0.0018198441],"genre_scores_gemma":[0.82336295,0.00034577108,0.16102278,0.00014239695,0.00013588402,0.00036510464,0.011400206,0.00012627937,0.003098714],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999403,0.000120344455,0.000055798293,0.00021963652,0.00012293813,0.00007831015],"domain_scores_gemma":[0.99898726,0.00038895514,0.00014764309,0.00016822711,0.00025553393,0.00005235593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010507376,0.00087297213,0.00065265974,0.0031383603,0.00022318133,0.00083931716,0.0010051008,0.0006137836,0.0010228655],"category_scores_gemma":[0.0026870985,0.00030111257,0.0014520143,0.0023831236,0.00018802927,0.0012726305,0.00086521154,0.00073004403,0.0012292835],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040846466,0.0006338488,0.09211027,0.00029400972,0.00039775396,0.0005172023,0.0003702302,0.27509424,0.018525997,0.005022388,0.013455787,0.59316975],"study_design_scores_gemma":[0.0000039272777,0.000029850338,0.004157341,0.000004934815,0.000018149227,0.0000342628,0.000029812785,0.9921565,0.0013756958,0.0012512414,0.0009286863,0.000009619899],"about_ca_topic_score_codex":0.006123237,"about_ca_topic_score_gemma":0.0095334565,"teacher_disagreement_score":0.006123237,"about_ca_system_score_codex":0.00047010975,"about_ca_system_score_gemma":0.0006389374,"threshold_uncertainty_score":0.012175202},"labels":[],"label_agreement":null},{"id":"W4319986073","doi":"10.17760/d20467253","title":"Self-repetition in abstractive neural summarizers","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"U.S. National Library of Medicine; National Institutes of Health; National Science Foundation","keywords":"Repetition (rhetorical device); Computer science; Natural language processing; Artificial intelligence; Qualitative analysis; Linguistics; Qualitative research; Sociology","score_opus":0.014088246249045555,"score_gpt":0.2281941035851638,"score_spread":0.21410585733611823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4319986073","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84463495,0.0016422723,0.1449842,0.00047993817,0.00009790593,0.00017787731,0.001581029,0.0022788097,0.0041229962],"genre_scores_gemma":[0.9659938,0.00019675378,0.030480392,0.000059337515,0.00004958908,0.00007928032,0.0016890374,0.00014599999,0.0013057436],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997273,0.00088758644,0.00032440393,0.00047893042,0.000883066,0.00015292705],"domain_scores_gemma":[0.96101165,0.024680682,0.005661342,0.0039531146,0.004199003,0.0004943152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005078892,0.00054510025,0.0006509031,0.0022142911,0.000501806,0.0014270786,0.0006494797,0.0005909236,0.001350554],"category_scores_gemma":[0.03568418,0.00023334021,0.00047282895,0.0013616068,0.00062045426,0.0022734033,0.00096624787,0.0008325906,0.00046613393],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018790624,0.00046910506,0.13596311,0.002284057,0.0011604185,0.0010516397,0.00685645,0.17641969,0.113895595,0.010120882,0.007655892,0.5422441],"study_design_scores_gemma":[0.000058207726,0.0019579711,0.13248311,0.00020597441,0.00042991602,0.0011830964,0.001701926,0.7365941,0.09592291,0.017989837,0.01127281,0.00020012734],"about_ca_topic_score_codex":0.00094137515,"about_ca_topic_score_gemma":0.001692201,"teacher_disagreement_score":0.005078892,"about_ca_system_score_codex":0.0005811927,"about_ca_system_score_gemma":0.0002830291,"threshold_uncertainty_score":0.026860118},"labels":[],"label_agreement":null},{"id":"W4320075798","doi":"10.15173/sciential.vi8.3014","title":"Health in Bite Sized Pieces - Discovering Lack of Accessibility and Engagement in Lay Summaries","year":2022,"lang":"en","type":"article","venue":"Sciential - McMaster Undergraduate Science Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Rubric; Psychology; Medicine; Medical education; Mathematics education","score_opus":0.055888154078025123,"score_gpt":0.3317369690351324,"score_spread":0.2758488149571073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4320075798","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9606711,0.0033116133,0.014102949,0.004283597,0.00038212357,0.0006916078,0.0012368117,0.00030901784,0.015011129],"genre_scores_gemma":[0.98604417,0.0011905354,0.008980747,0.00066161394,0.00022713676,0.00042825183,0.0004973159,0.00006715366,0.0019031616],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.967208,0.018894278,0.0055071986,0.0013737556,0.0062079895,0.0008087469],"domain_scores_gemma":[0.69288856,0.18014443,0.07625819,0.016985416,0.027353916,0.0063694054],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.029988384,0.0004194843,0.0005611853,0.005078042,0.0013488374,0.0051244656,0.00068603054,0.00073035614,0.0058198166],"category_scores_gemma":[0.2644218,0.0003878673,0.0007114774,0.0026486963,0.0014297611,0.005293311,0.0053240326,0.0011182401,0.00088274427],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010792105,0.00036571038,0.37193748,0.0047150524,0.0005237233,0.0014190786,0.25199354,0.00046212957,0.004743003,0.005491665,0.015691696,0.3415777],"study_design_scores_gemma":[0.00013116233,0.002092024,0.622011,0.0043656095,0.00054805045,0.004146787,0.23550467,0.002595297,0.0045402385,0.013969517,0.10980031,0.0002953399],"about_ca_topic_score_codex":0.000564889,"about_ca_topic_score_gemma":0.001059589,"teacher_disagreement_score":0.99487555,"about_ca_system_score_codex":0.0011998843,"about_ca_system_score_gemma":0.0013780954,"threshold_uncertainty_score":0.15859562},"labels":[],"label_agreement":null},{"id":"W4320853823","doi":"10.48550/arxiv.2302.05895","title":"Discourse Structure Extraction from Pre-Trained and Fine-Tuned Language Models in Dialogues","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Agence Nationale de la Recherche","keywords":"Computer science; Task (project management); Exploit; Natural language processing; Artificial intelligence; Sentence; Language model; Machine learning","score_opus":0.08395320489432903,"score_gpt":0.22297286794082338,"score_spread":0.13901966304649435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4320853823","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13432387,0.0018697671,0.84260666,0.00076370424,0.00021434481,0.00026740463,0.002198236,0.014722118,0.0030338492],"genre_scores_gemma":[0.6327156,0.0005958906,0.35291114,0.00017184872,0.0002103666,0.00041222206,0.007987593,0.0010412358,0.0039541437],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99724066,0.0014309495,0.00012975333,0.0008320909,0.00020875018,0.00015787975],"domain_scores_gemma":[0.99445266,0.0038428586,0.00024530495,0.0005570119,0.00072049233,0.00018165514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026816218,0.0018178471,0.0010711526,0.0026313611,0.0008976506,0.0019767908,0.0013042942,0.0014300102,0.002178216],"category_scores_gemma":[0.011642795,0.00091235165,0.0013289701,0.00136393,0.0006311833,0.0028737944,0.0018087167,0.0026035295,0.0026993451],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010457371,0.00047812675,0.0071298857,0.0011510126,0.0003797693,0.0004674388,0.003309898,0.12182213,0.08897983,0.0063776714,0.016905546,0.75195307],"study_design_scores_gemma":[0.000081990445,0.00017087771,0.0030219238,0.00007784577,0.00011399063,0.00010154445,0.0006856611,0.94873977,0.028150205,0.009907348,0.008894955,0.000053881988],"about_ca_topic_score_codex":0.0046861665,"about_ca_topic_score_gemma":0.009444249,"teacher_disagreement_score":0.0046861665,"about_ca_system_score_codex":0.0011479689,"about_ca_system_score_gemma":0.0015226173,"threshold_uncertainty_score":0.014181972},"labels":[],"label_agreement":null},{"id":"W4321012596","doi":"10.18653/v1/2023.eacl-main.167","title":"Exploring Category Structure with Contextual Language Models and Lexical Semantic Networks","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Natural language processing; Artificial intelligence; WordNet; Polysemy; Similarity (geometry); Task (project management); Word (group theory); Semantic similarity; Priming (agriculture); Linguistics","score_opus":0.12428779778739875,"score_gpt":0.2696044762852646,"score_spread":0.1453166784978659,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321012596","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38646457,0.0015715332,0.603706,0.00080562156,0.00008845301,0.00006088396,0.001566519,0.0012313215,0.0045051104],"genre_scores_gemma":[0.9301559,0.0003486446,0.0667959,0.00010285001,0.00006258727,0.000072832394,0.0014064395,0.00009958769,0.0009552088],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994123,0.0002852988,0.000025892668,0.0001860514,0.000046629793,0.000043917826],"domain_scores_gemma":[0.99661666,0.0023955365,0.00027656532,0.00036043546,0.00021090887,0.00013990715],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014322188,0.0009565964,0.0006367053,0.0023921838,0.0004607743,0.0015982145,0.000756643,0.00087685103,0.002587802],"category_scores_gemma":[0.007405586,0.00042709714,0.0009460368,0.0018347778,0.00059983024,0.0052726534,0.0012226613,0.0014049741,0.0008988061],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011373707,0.00056483346,0.0877152,0.00070361595,0.0007363024,0.00044458002,0.0019201929,0.37703657,0.01421283,0.09454048,0.00817838,0.41280955],"study_design_scores_gemma":[0.000015269548,0.000061993014,0.0040207948,0.000030458483,0.000043231816,0.00008839929,0.00016710427,0.85396725,0.0008355051,0.13905537,0.0016915775,0.000023043252],"about_ca_topic_score_codex":0.003715665,"about_ca_topic_score_gemma":0.007636516,"teacher_disagreement_score":0.003715665,"about_ca_system_score_codex":0.0007204744,"about_ca_system_score_gemma":0.0005795167,"threshold_uncertainty_score":0.008657038},"labels":[],"label_agreement":null},{"id":"W4321175651","doi":"10.48550/arxiv.2302.07738","title":"Alloprof: a new French question-answer education dataset and its use in an information retrieval case study","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Question answering; Relevance (law); Computer science; Context (archaeology); Spelling; Information retrieval; Task (project management); Variety (cybernetics); Comprehension; Baseline (sea); Natural language processing; Artificial intelligence; World Wide Web; Linguistics","score_opus":0.15646335249567606,"score_gpt":0.25496705731576963,"score_spread":0.09850370482009357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321175651","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13283259,0.0034647067,0.011548076,0.0023911255,0.00042290703,0.0017572427,0.8168277,0.014811121,0.01594452],"genre_scores_gemma":[0.041948803,0.00032456912,0.020675305,0.00043040863,0.00008315568,0.00095248676,0.9311487,0.00028329733,0.004153257],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99722785,0.00073788903,0.0003170253,0.0007061069,0.0007054123,0.0003056386],"domain_scores_gemma":[0.9952153,0.0012393976,0.0002192663,0.00081671827,0.0020694654,0.00043989162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020658998,0.0020834478,0.00093827647,0.0063878573,0.0018226069,0.0017327153,0.002410339,0.0033675598,0.0080956565],"category_scores_gemma":[0.008170144,0.00030319474,0.0014680449,0.00453448,0.00064492016,0.0018200236,0.001824711,0.0017448091,0.006829577],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009002752,0.0018421103,0.020259412,0.0034359286,0.00022000555,0.0014706345,0.0015270554,0.0061581456,0.007156554,0.0031507744,0.84442234,0.10945679],"study_design_scores_gemma":[0.0005430167,0.00058478286,0.0705695,0.0006217262,0.00013035147,0.0015627356,0.0024462286,0.031266168,0.00978073,0.001923705,0.8803284,0.00024267731],"about_ca_topic_score_codex":0.17684403,"about_ca_topic_score_gemma":0.26076075,"teacher_disagreement_score":0.17684403,"about_ca_system_score_codex":0.004073305,"about_ca_system_score_gemma":0.0026752367,"threshold_uncertainty_score":0.35162938},"labels":[],"label_agreement":null},{"id":"W4321471877","doi":"","title":"Identifying Similar Test Cases That Are Specified in Natural Language","year":2022,"lang":"en","type":"report","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Test (biology); Natural (archaeology); Computer science; Natural language processing; Linguistics; Artificial intelligence; Geography; Geology; Archaeology; Philosophy; Paleontology","score_opus":0.049587667605116806,"score_gpt":0.27667296410581743,"score_spread":0.22708529650070064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4321471877","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29336983,0.0009247142,0.681003,0.0010936615,0.00030837613,0.0024782037,0.0050411765,0.009353504,0.006427533],"genre_scores_gemma":[0.55976737,0.00046848896,0.40629336,0.0009569578,0.00020633053,0.0024761816,0.025259737,0.001628239,0.0029433956],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9840552,0.004386422,0.002103593,0.003655202,0.004948262,0.00085131],"domain_scores_gemma":[0.9404731,0.04206364,0.0057782675,0.005301013,0.0057349335,0.000648977],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004974076,0.0021100794,0.0014348496,0.008106047,0.0010501867,0.002469492,0.003223129,0.002318822,0.003159827],"category_scores_gemma":[0.055559635,0.00063332997,0.0030306473,0.004144728,0.001764443,0.0033823107,0.0017939013,0.0017325974,0.0011177331],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018663233,0.0033864076,0.11298563,0.004981099,0.0013092472,0.024745194,0.006050654,0.06852412,0.09079995,0.055522393,0.0316599,0.59816915],"study_design_scores_gemma":[0.00066178065,0.0015544397,0.044233464,0.0011222698,0.00095688924,0.01917311,0.0040331786,0.6623822,0.08611485,0.102424815,0.07691582,0.00042718684],"about_ca_topic_score_codex":0.0064395606,"about_ca_topic_score_gemma":0.0069746855,"teacher_disagreement_score":0.008106047,"about_ca_system_score_codex":0.0018592089,"about_ca_system_score_gemma":0.0031706565,"threshold_uncertainty_score":0.026305795},"labels":[],"label_agreement":null},{"id":"W4322096525","doi":"10.1007/978-3-031-24337-0_22","title":"Building Personalized Language Models Through Language Model Interpolation","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Perplexity; Language model; Computer science; Artificial intelligence; Natural language processing; Social media; Interpolation (computer graphics); Cache language model; Natural language; Universal Networking Language; Comprehension approach; World Wide Web","score_opus":0.0376111306386286,"score_gpt":0.2919027307924812,"score_spread":0.2542916001538526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4322096525","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004134963,0.00021484865,0.98473686,0.000093689254,0.000056195247,0.00005458496,0.000551586,0.008385957,0.0017713],"genre_scores_gemma":[0.14119925,0.00089092227,0.83607334,0.0002511814,0.000114603,0.00031542932,0.0057645766,0.0030307868,0.012359891],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99898154,0.0002734562,0.000052495827,0.00035291153,0.00025445348,0.00008507127],"domain_scores_gemma":[0.9984932,0.0008202241,0.000052617284,0.00039854622,0.00018346892,0.00005201856],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011762629,0.0012322735,0.0012140452,0.0013683139,0.0007293257,0.001995872,0.0018015605,0.0011774959,0.009409558],"category_scores_gemma":[0.004298896,0.0015965612,0.00293874,0.0015884852,0.0004902983,0.003971796,0.0019405725,0.0030230628,0.009681283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036036805,0.00027049714,0.001399156,0.000362267,0.0003522592,0.00037849427,0.0005208606,0.23983498,0.018230755,0.030739289,0.02958912,0.67796195],"study_design_scores_gemma":[0.00002181518,0.000037670896,0.00017962852,0.000021307349,0.00008154415,0.00014914837,0.000063347754,0.94999987,0.0084454045,0.030150436,0.010820147,0.000029756768],"about_ca_topic_score_codex":0.006201977,"about_ca_topic_score_gemma":0.010743846,"teacher_disagreement_score":0.009409558,"about_ca_system_score_codex":0.00085193414,"about_ca_system_score_gemma":0.0010677151,"threshold_uncertainty_score":0.031478047},"labels":[],"label_agreement":null},{"id":"W4323027187","doi":"10.2139/ssrn.4369067","title":"Explanation-Oriented (vs. Free) Discussion Promotes Open-Minded Political Reasoning","year":2023,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Politics; Political science; Epistemology; Motivated reasoning; Computer science; Positive economics; Economics; Law; Philosophy","score_opus":0.015338100956494637,"score_gpt":0.27751531713419986,"score_spread":0.26217721617770523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323027187","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5957711,0.000556604,0.17162071,0.013442388,0.0002888027,0.00040776038,0.00023782604,0.0007178414,0.216957],"genre_scores_gemma":[0.98787874,0.00009594993,0.008752152,0.0003196206,0.00008114715,0.00006371742,0.00006473139,0.00007397457,0.0026699647],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.989588,0.0071607204,0.00029875155,0.0010108622,0.0011784332,0.00076330366],"domain_scores_gemma":[0.90185523,0.08164652,0.005733852,0.0056637116,0.0022466276,0.0028541915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014699254,0.00040376838,0.00048218187,0.0015628574,0.0018129409,0.005435018,0.0010821956,0.003356998,0.020004408],"category_scores_gemma":[0.05321526,0.00034601882,0.00070327945,0.0009426165,0.0031485658,0.007936072,0.0050557125,0.0030155901,0.0014728879],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002096051,0.001922225,0.03945808,0.00096226396,0.00018421999,0.0005667536,0.046337962,0.007267631,0.012612508,0.72814274,0.009667039,0.1507825],"study_design_scores_gemma":[0.0005429508,0.00035785866,0.020214563,0.00017444615,0.00021444047,0.0002697974,0.010636576,0.02785667,0.00417653,0.90548366,0.030005634,0.000067007924],"about_ca_topic_score_codex":0.0009379994,"about_ca_topic_score_gemma":0.0013361947,"teacher_disagreement_score":0.020004408,"about_ca_system_score_codex":0.0015039708,"about_ca_system_score_gemma":0.001799078,"threshold_uncertainty_score":0.07773799},"labels":[],"label_agreement":null},{"id":"W4323076538","doi":"10.48550/arxiv.2303.01410","title":"NLP Workbench: Efficient and Extensible Integration of State-of-the-art Text Mining Tools","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Workbench; Computer science; Parsing; Extensibility; License; MIT License; Artificial intelligence; Architecture; Sentiment analysis; Natural language processing; Information retrieval; Visualization; Programming language","score_opus":0.13289432239835935,"score_gpt":0.2043208972299189,"score_spread":0.07142657483155954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323076538","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007475762,0.0005222841,0.63255155,0.0005570353,0.0003084802,0.000864651,0.013750736,0.33524144,0.0087279985],"genre_scores_gemma":[0.074809074,0.0011763177,0.7574185,0.0007486045,0.00028355245,0.002144431,0.11186466,0.03393842,0.017616384],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979785,0.00031892877,0.00027924954,0.00060757564,0.00068813557,0.00012762804],"domain_scores_gemma":[0.99661416,0.0015923256,0.00017665935,0.00079163717,0.0005727214,0.0002524377],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028843444,0.0022629953,0.0012439217,0.004837975,0.0010456762,0.0033714687,0.003412203,0.0012378019,0.01484363],"category_scores_gemma":[0.008679477,0.0013538941,0.0015005656,0.0033914559,0.0008516929,0.0062300656,0.0043968847,0.0019468988,0.013254991],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011131902,0.0005914471,0.0039423173,0.0014439326,0.00041980378,0.0018200732,0.0015285217,0.010710711,0.03384821,0.015279856,0.37507707,0.5542249],"study_design_scores_gemma":[0.00055345346,0.00032530277,0.003314633,0.0003462173,0.00018412572,0.0011427029,0.00076931453,0.38470572,0.07255174,0.048443217,0.48735502,0.0003085514],"about_ca_topic_score_codex":0.00423219,"about_ca_topic_score_gemma":0.004850595,"teacher_disagreement_score":0.01484363,"about_ca_system_score_codex":0.0008233242,"about_ca_system_score_gemma":0.0022481254,"threshold_uncertainty_score":0.049656928},"labels":[],"label_agreement":null},{"id":"W4323262652","doi":"10.1016/j.cogsys.2023.02.009","title":"Computing word meanings by aggregating individualized distributional models: Wisdom of the crowds in lexical semantic memory","year":2023,"lang":"en","type":"article","venue":"Cognitive Systems Research","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Similarity (geometry); Judgement; Semantic similarity; Word (group theory); Computer science; Semantic memory; Natural language processing; Aggregate (composite); Artificial intelligence; Psychology; Cognitive psychology; Cognition; Linguistics","score_opus":0.13329042708033806,"score_gpt":0.3754494462827638,"score_spread":0.24215901920242572,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323262652","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31490126,0.002582781,0.67364734,0.0047159996,0.00028283425,0.00010054866,0.00036799975,0.00044101072,0.0029601525],"genre_scores_gemma":[0.91786754,0.0007705017,0.07903577,0.00037419266,0.0004158864,0.000102643346,0.00040711573,0.0001364333,0.0008898452],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99477094,0.002961324,0.00040885602,0.0011179486,0.00053841865,0.00020256429],"domain_scores_gemma":[0.92199343,0.06722436,0.002130692,0.0055495156,0.0021321413,0.000969867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01496651,0.0011605434,0.00314056,0.0040307087,0.001726616,0.006398574,0.0027147764,0.0032501733,0.0019094561],"category_scores_gemma":[0.1003358,0.0011919627,0.001658943,0.0042202906,0.0030648324,0.017423047,0.0044049746,0.004116752,0.000491754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011774172,0.00068516197,0.039602056,0.000632804,0.0014406996,0.00045342618,0.006133634,0.3699652,0.0020524957,0.18611951,0.009301227,0.38243636],"study_design_scores_gemma":[0.000032927466,0.000035866273,0.0015523861,0.000032202854,0.000066799395,0.000046700206,0.00042295025,0.51246053,0.0002967888,0.48443192,0.0005816709,0.000039285034],"about_ca_topic_score_codex":0.006718638,"about_ca_topic_score_gemma":0.007835204,"teacher_disagreement_score":0.01496651,"about_ca_system_score_codex":0.001484164,"about_ca_system_score_gemma":0.0014692366,"threshold_uncertainty_score":0.07915139},"labels":[],"label_agreement":null},{"id":"W4323649518","doi":"10.23977/acss.2023.070111","title":"Research on Named Entity Recognition Method Based on Language Pre-Training Model","year":2023,"lang":"en","type":"article","venue":"Advances in Computer Signals and Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Word (group theory); Representation (politics); Polysemy; Artificial intelligence; Word embedding; Set (abstract data type); Vectorization (mathematics); Natural language processing; Embedding; Language model; Training set; Support vector machine; Data set; Machine learning; Data mining; Mathematics","score_opus":0.16248387250087376,"score_gpt":0.423410809705627,"score_spread":0.26092693720475324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323649518","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0109537225,0.0011645602,0.98160833,0.00028745356,0.00016738915,0.000079108715,0.00027465474,0.0039509586,0.0015137458],"genre_scores_gemma":[0.3668886,0.0046264133,0.6049464,0.00045591008,0.00035411373,0.00040179817,0.0056478763,0.00062323094,0.016055675],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982551,0.00037615836,0.00016124982,0.00066847925,0.00040524284,0.00013385034],"domain_scores_gemma":[0.9984156,0.00046763627,0.00011996739,0.0003683573,0.0005727677,0.000055735884],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014517889,0.0012730605,0.0011644185,0.001709233,0.00055505836,0.0013779746,0.0018029569,0.0008309021,0.003209818],"category_scores_gemma":[0.004168211,0.00042741932,0.0011898244,0.0022702082,0.00049348857,0.007165383,0.0009278521,0.0019045669,0.0029799812],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016730536,0.00014420177,0.0023560717,0.00028163684,0.00012021765,0.00011883419,0.00020752521,0.031178579,0.016210513,0.009124467,0.007961,0.9321295],"study_design_scores_gemma":[0.000028329263,0.00017845561,0.0024134247,0.000050863528,0.00010817409,0.0003672495,0.00023076347,0.928447,0.04309628,0.008050895,0.016944176,0.00008440047],"about_ca_topic_score_codex":0.005507895,"about_ca_topic_score_gemma":0.0029283867,"teacher_disagreement_score":0.005507895,"about_ca_system_score_codex":0.0006428411,"about_ca_system_score_gemma":0.0011228652,"threshold_uncertainty_score":0.010951638},"labels":[],"label_agreement":null},{"id":"W4323663117","doi":"10.1007/s00521-023-08393-4","title":"DeepPress: guided press release topic-aware text generation using ensemble transformers","year":2023,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Transformer; Fluency; Artificial intelligence; Continuation; Variety (cybernetics); Key (lock); Context (archaeology); Focus (optics); Natural language processing; Writing style; Linguistics","score_opus":0.08579680199765757,"score_gpt":0.3175017096316356,"score_spread":0.23170490763397802,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323663117","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012993478,0.0007665856,0.84130245,0.0019995626,0.002736115,0.0004948297,0.015162576,0.09751487,0.027029524],"genre_scores_gemma":[0.28037366,0.0012321094,0.5501826,0.0008814386,0.0017030357,0.0009972437,0.059148964,0.019750493,0.08573042],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99894303,0.00021280762,0.00007587231,0.00029137245,0.00035289943,0.00012404473],"domain_scores_gemma":[0.9971033,0.0009326337,0.00008817784,0.0008268486,0.0008501645,0.0001989278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015876688,0.0011835712,0.000840826,0.001533723,0.0005751221,0.0023393843,0.0013565458,0.0013900183,0.07254812],"category_scores_gemma":[0.008718103,0.0006683675,0.0012548288,0.0015552723,0.0003808075,0.00409265,0.002328869,0.0023373475,0.040248495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065898075,0.00023327087,0.0005408042,0.0003569406,0.00009411468,0.00016314279,0.00012470143,0.017281193,0.019522648,0.022801258,0.3158459,0.62237704],"study_design_scores_gemma":[0.00028758805,0.00026850365,0.000868637,0.0000844082,0.000081642414,0.00029682266,0.00008681198,0.7022584,0.05279585,0.08245674,0.16043025,0.00008435236],"about_ca_topic_score_codex":0.0016833456,"about_ca_topic_score_gemma":0.0023945908,"teacher_disagreement_score":0.07254812,"about_ca_system_score_codex":0.0005794449,"about_ca_system_score_gemma":0.0013118453,"threshold_uncertainty_score":0.24269766},"labels":[],"label_agreement":null},{"id":"W4323782875","doi":"10.1007/978-3-031-25891-6_23","title":"A Comparison of SVM Against Pre-trained Language Models (PLMs) for Text Classification Tasks","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Western University","funders":"","keywords":"Computer science; Support vector machine; Artificial intelligence; Classifier (UML); Machine learning; Feature engineering; Natural language processing; Task (project management); Feature (linguistics); Linear classifier; Pattern recognition (psychology); Deep learning","score_opus":0.06515733099742363,"score_gpt":0.3215682359011537,"score_spread":0.2564109049037301,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323782875","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56813174,0.030056156,0.3494113,0.0016890262,0.0039743385,0.00043127587,0.0055118115,0.02691539,0.013879031],"genre_scores_gemma":[0.8318294,0.0033636838,0.13550776,0.00041536425,0.0005011343,0.00023339487,0.014141508,0.00087192125,0.013135683],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981267,0.00072597555,0.00018842635,0.00039384555,0.0004135226,0.00015154046],"domain_scores_gemma":[0.99286807,0.004679858,0.00015329548,0.0004887365,0.0015647908,0.00024532177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037584887,0.0014091013,0.0013504319,0.0013213841,0.0005464577,0.0015070065,0.0012568759,0.0014097193,0.003510776],"category_scores_gemma":[0.0073814895,0.00027708412,0.00076385075,0.0012572836,0.00020465672,0.0022804055,0.0009397625,0.0016190683,0.0030509327],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004330931,0.00072468776,0.004748824,0.00057196897,0.00054663554,0.00012140097,0.00010366254,0.040098608,0.011883518,0.0007040433,0.020515976,0.91564983],"study_design_scores_gemma":[0.00013533865,0.0011493017,0.0049699266,0.00005897618,0.00018390192,0.00014196671,0.0001600586,0.97776216,0.010369652,0.0013912262,0.0036472727,0.000030148896],"about_ca_topic_score_codex":0.005590548,"about_ca_topic_score_gemma":0.0050890367,"teacher_disagreement_score":0.005590548,"about_ca_system_score_codex":0.00082346925,"about_ca_system_score_gemma":0.00088857504,"threshold_uncertainty_score":0.019877017},"labels":[],"label_agreement":null},{"id":"W4324355422","doi":"10.3390/data8030061","title":"TKGQA Dataset: Using Question Answering to Guide and Validate the Evolution of Temporal Knowledge Graph","year":2023,"lang":"en","type":"article","venue":"Data","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Bank of Canada","funders":"","keywords":"Knowledge graph; Computer science; Question answering; Graph; Parsing; Information retrieval; Process (computing); Artificial intelligence; Natural language processing; Theoretical computer science","score_opus":0.1071454944889778,"score_gpt":0.3625074726673969,"score_spread":0.2553619781784191,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4324355422","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033454396,0.0020732814,0.009118727,0.0013053208,0.0003325365,0.00046740807,0.9327068,0.014067033,0.0064745573],"genre_scores_gemma":[0.027244126,0.00029161066,0.018168576,0.00021672483,0.000048001977,0.0002329199,0.95188713,0.00019653594,0.0017143788],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979534,0.00042860673,0.00032076004,0.00066723063,0.0004852306,0.00014474039],"domain_scores_gemma":[0.9947412,0.0018366505,0.0005363427,0.00116533,0.001361989,0.0003583642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016326498,0.002008792,0.0007875399,0.0055491356,0.0013318051,0.0018265959,0.0024305105,0.0029468995,0.0060471264],"category_scores_gemma":[0.010199364,0.00038179022,0.0015853415,0.004287284,0.00063718064,0.0034247215,0.0016188378,0.0018968045,0.007039062],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008109131,0.0007023546,0.017783405,0.0037288824,0.00038990722,0.00080881425,0.001001113,0.0081206355,0.008575771,0.0045009055,0.8671174,0.08645992],"study_design_scores_gemma":[0.00056612934,0.00039969507,0.04934801,0.0006039458,0.00035222273,0.0011509554,0.002123306,0.0876628,0.0137598105,0.01055283,0.8332495,0.00023082463],"about_ca_topic_score_codex":0.054395188,"about_ca_topic_score_gemma":0.0856514,"teacher_disagreement_score":0.054395188,"about_ca_system_score_codex":0.002205411,"about_ca_system_score_gemma":0.0019228589,"threshold_uncertainty_score":0.10815716},"labels":[],"label_agreement":null},{"id":"W4327498258","doi":"10.1007/978-3-031-28241-6_38","title":"Uncertainty Quantification for Text Classification","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Uncertainty quantification; Artificial intelligence; Robustness (evolution); Machine learning; Scalability; Generalization; Context (archaeology); Dropout (neural networks); Bayesian probability; Ensemble learning; Deep learning; Data mining; Mathematics; Database","score_opus":0.07459716417679207,"score_gpt":0.2958613792212944,"score_spread":0.22126421504450233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327498258","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031205702,0.0045636515,0.9871568,0.00066037435,0.00016078439,0.000048496055,0.00036128354,0.000497543,0.0034303826],"genre_scores_gemma":[0.36100456,0.007505343,0.60341537,0.0008412558,0.0020560657,0.0006913953,0.0035986209,0.0007250349,0.020162422],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963659,0.0012533254,0.00029405183,0.00060974114,0.0012786643,0.00019823376],"domain_scores_gemma":[0.9920993,0.0059487615,0.00040199142,0.0007092522,0.0006818978,0.00015884754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042707054,0.0011615411,0.001959741,0.003469139,0.0010571807,0.003400495,0.0020790643,0.0016439039,0.0050652656],"category_scores_gemma":[0.0147172,0.0007424428,0.001962192,0.003397552,0.0017446473,0.0064561847,0.0030594696,0.0036257005,0.0015158982],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014864476,0.00008301225,0.00060441735,0.0006761404,0.0001880055,0.00009359489,0.00028726476,0.08934406,0.003339403,0.33291793,0.023475075,0.5488425],"study_design_scores_gemma":[0.0000059318395,0.000029591898,0.00037779394,0.00010421486,0.000032442556,0.000053942967,0.000046541893,0.45327282,0.0013542145,0.53863144,0.0060622883,0.000028767248],"about_ca_topic_score_codex":0.0025696848,"about_ca_topic_score_gemma":0.001788627,"teacher_disagreement_score":0.0050652656,"about_ca_system_score_codex":0.0023426283,"about_ca_system_score_gemma":0.0010064524,"threshold_uncertainty_score":0.022585928},"labels":[],"label_agreement":null},{"id":"W4327498763","doi":"10.1007/978-3-031-28241-6_10","title":"PyGaggle: A Gaggle of Resources for Open-Domain Question Answering","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Question answering; Codebase; Information retrieval; Open domain; Artificial intelligence; Context (archaeology); Domain (mathematical analysis); Machine learning; Programming language; Source code","score_opus":0.03306528010206973,"score_gpt":0.2818557621392487,"score_spread":0.24879048203717896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327498763","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032488003,0.0011623183,0.6911969,0.0011653317,0.0003139036,0.00040258368,0.009684407,0.26894546,0.023880303],"genre_scores_gemma":[0.07584271,0.0017327073,0.7821374,0.0016878551,0.0002553659,0.0014279963,0.047889676,0.05307853,0.035947736],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9980233,0.0007186595,0.00016379812,0.0004096519,0.00051701046,0.00016760669],"domain_scores_gemma":[0.99576,0.002458624,0.00008533421,0.0011248395,0.00027197634,0.00029926133],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022769393,0.0017831459,0.0014903516,0.0032052533,0.0011802056,0.003967965,0.0032118917,0.0021796622,0.05273448],"category_scores_gemma":[0.010810399,0.0015301076,0.001966507,0.0031180247,0.0014474557,0.0092970645,0.0100217,0.003685567,0.037805483],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086990517,0.00019075794,0.00083295815,0.0011859392,0.000106604704,0.00043337152,0.0015412646,0.0049531893,0.007441981,0.07591446,0.46715838,0.4393712],"study_design_scores_gemma":[0.00019066726,0.00007567026,0.0004490098,0.0003555683,0.0000701217,0.0004976166,0.00051012775,0.050533302,0.011503331,0.19405146,0.7416366,0.00012658552],"about_ca_topic_score_codex":0.0038382025,"about_ca_topic_score_gemma":0.0053593875,"teacher_disagreement_score":0.05273448,"about_ca_system_score_codex":0.0010685745,"about_ca_system_score_gemma":0.0016662206,"threshold_uncertainty_score":0.17641443},"labels":[],"label_agreement":null},{"id":"W4327498768","doi":"10.1007/978-3-031-28241-6_11","title":"Pre-processing Matters! Improved Wikipedia Corpora for Open-Domain Question Answering","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Question answering; Natural language processing; Information retrieval; Domain (mathematical analysis); Landmark; Artificial intelligence; Text corpus; Open domain; Pipeline (software)","score_opus":0.028800474068406422,"score_gpt":0.27818040762012664,"score_spread":0.24937993355172022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327498768","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.089596,0.0074307043,0.31039405,0.004558771,0.005320698,0.001732398,0.40187454,0.11182419,0.06726865],"genre_scores_gemma":[0.08103187,0.001245496,0.24851751,0.00053502765,0.0007862792,0.0012432489,0.6391909,0.007877875,0.019571815],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99730116,0.00084959227,0.00030739838,0.00073998776,0.0006245174,0.00017730014],"domain_scores_gemma":[0.98304224,0.006757754,0.00046446844,0.0029644023,0.006096799,0.00067437097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026105212,0.001449152,0.0012388504,0.005171527,0.0017116137,0.0028695113,0.0016759951,0.0012776356,0.028227212],"category_scores_gemma":[0.017618349,0.0010528989,0.001060559,0.0042210277,0.00056884886,0.0051614936,0.0031876825,0.0028654863,0.028439747],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040080826,0.0004847806,0.0027938425,0.0022279203,0.00018267732,0.00040079473,0.0010552914,0.0029580353,0.028920325,0.0067014983,0.63796824,0.31590578],"study_design_scores_gemma":[0.00026654816,0.00023631868,0.012168809,0.0005128859,0.00031810344,0.00072266633,0.0016059817,0.05759106,0.05819633,0.017771108,0.85037506,0.00023521132],"about_ca_topic_score_codex":0.006529718,"about_ca_topic_score_gemma":0.014808228,"teacher_disagreement_score":0.028227212,"about_ca_system_score_codex":0.00064519007,"about_ca_system_score_gemma":0.002021072,"threshold_uncertainty_score":0.09442943},"labels":[],"label_agreement":null},{"id":"W4327644054","doi":"10.1007/978-3-031-28238-6_28","title":"Adversarial Adaptation for French Named Entity Recognition","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Overfitting; Domain adaptation; Transformer; Named-entity recognition; Artificial intelligence; Adversarial system; Natural language processing; Machine learning; Task (project management); Adaptation (eye); Language model; Artificial neural network; Classifier (UML)","score_opus":0.05240344038885753,"score_gpt":0.25863514913071334,"score_spread":0.2062317087418558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644054","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010104534,0.001816216,0.9748955,0.0005264513,0.00037836717,0.000055280398,0.00073179515,0.0072401725,0.0042517805],"genre_scores_gemma":[0.47133166,0.0025769244,0.46157762,0.0008601454,0.00063798414,0.00032796114,0.009972571,0.002138801,0.05057635],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989561,0.00045446248,0.000042197888,0.0002868184,0.00014091983,0.000119578755],"domain_scores_gemma":[0.9987514,0.0007504355,0.00004436719,0.00024494354,0.00017540838,0.000033417688],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018025477,0.0014030902,0.0010552071,0.0008463923,0.0006083475,0.0011680311,0.0016502214,0.0014791031,0.0056023593],"category_scores_gemma":[0.0030680439,0.00059169,0.0012760406,0.001277907,0.0005596333,0.0015340457,0.0017681909,0.0022620654,0.0050588027],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028668586,0.000084235784,0.00057419145,0.00012989709,0.00018662345,0.0002177021,0.00009949189,0.33153072,0.008156288,0.014396293,0.038198866,0.606139],"study_design_scores_gemma":[0.0000063454304,0.000021755068,0.00027886298,0.000013912798,0.000019464533,0.00006198642,0.000016080685,0.9849902,0.0036649855,0.006091805,0.004820428,0.000014105179],"about_ca_topic_score_codex":0.012543697,"about_ca_topic_score_gemma":0.012595857,"teacher_disagreement_score":0.012543697,"about_ca_system_score_codex":0.00094876275,"about_ca_system_score_gemma":0.0006159058,"threshold_uncertainty_score":0.024941385},"labels":[],"label_agreement":null},{"id":"W4327644059","doi":"10.1007/978-3-031-28238-6_7","title":"Improving the Generalizability of the Dense Passage Retriever Using Generated Datasets","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Overfitting; Generalizability theory; Domain (mathematical analysis); Context (archaeology); Artificial intelligence; Information retrieval; Focus (optics); Encoder; Labrador Retriever; Training set; Data mining; Pattern recognition (psychology); Mathematics; Statistics","score_opus":0.042814254776952124,"score_gpt":0.25934468536600996,"score_spread":0.21653043058905785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644059","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30840647,0.008775915,0.45545334,0.0024160976,0.0017385742,0.001316998,0.045695644,0.16273496,0.013462038],"genre_scores_gemma":[0.47293293,0.002288706,0.3493857,0.001216209,0.0010240481,0.0007010159,0.15581752,0.006100635,0.010533138],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9931049,0.0021543347,0.00061954436,0.0022345635,0.0015101425,0.0003765755],"domain_scores_gemma":[0.9745285,0.012766729,0.00036035737,0.008650783,0.0033669886,0.00032655703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008357726,0.0024282467,0.0030262475,0.006605639,0.0012131771,0.0038484603,0.0033476606,0.002840825,0.006595774],"category_scores_gemma":[0.046749514,0.0008399564,0.0020849183,0.005525048,0.00071149034,0.006391327,0.0032162212,0.0023669165,0.010474981],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028181493,0.0012052713,0.013920643,0.0016666271,0.0015857914,0.000895124,0.0007304747,0.05054885,0.042143434,0.0037791363,0.14578404,0.73492247],"study_design_scores_gemma":[0.0007301657,0.0006285735,0.010108669,0.0001269472,0.00061607145,0.00088597275,0.0007103024,0.91729593,0.026643613,0.011295177,0.03081142,0.0001472439],"about_ca_topic_score_codex":0.017135173,"about_ca_topic_score_gemma":0.020140368,"teacher_disagreement_score":0.017135173,"about_ca_system_score_codex":0.00080575224,"about_ca_system_score_gemma":0.0020327372,"threshold_uncertainty_score":0.04420036},"labels":[],"label_agreement":null},{"id":"W4327644099","doi":"10.1007/978-3-031-28244-7_34","title":"CoSPLADE: Contextualizing SPLADE for Conversational Information Retrieval","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Agence Nationale de la Recherche","keywords":"Computer science; Ranking (information retrieval); Information retrieval; Task (project management); Conversation; Code (set theory); Rank (graph theory); Artificial intelligence; Source code; Programming language","score_opus":0.034249620816944965,"score_gpt":0.26238429948976755,"score_spread":0.22813467867282258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644099","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0055418382,0.00079278875,0.90842694,0.00016332776,0.0001991166,0.00029745576,0.0023792342,0.07619519,0.0060041975],"genre_scores_gemma":[0.07216174,0.00065445696,0.89242667,0.00031643038,0.00018821527,0.0006311987,0.010941357,0.00696944,0.01571058],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9991892,0.0002554891,0.00006737152,0.00020515366,0.00022418893,0.000058589856],"domain_scores_gemma":[0.9991543,0.00039605596,0.000028146067,0.00021343786,0.0001533012,0.000054723132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00093623245,0.0015516292,0.001083179,0.0013482614,0.00074318214,0.0016268124,0.0012926678,0.00083028234,0.02890184],"category_scores_gemma":[0.0029704592,0.000756868,0.0009260945,0.0010229435,0.0004562835,0.0025809703,0.003283612,0.0013260059,0.014183935],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001065377,0.00022825015,0.00062346586,0.0009856116,0.0001099797,0.00020581784,0.00094796496,0.0051380564,0.031664867,0.022947773,0.123849064,0.81223375],"study_design_scores_gemma":[0.00034742782,0.00039104332,0.0011959028,0.00023961811,0.00021231397,0.0006274981,0.0007424258,0.373152,0.083681256,0.068888046,0.47037327,0.00014921495],"about_ca_topic_score_codex":0.0025527668,"about_ca_topic_score_gemma":0.0049707307,"teacher_disagreement_score":0.02890184,"about_ca_system_score_codex":0.00044374037,"about_ca_system_score_gemma":0.0007277608,"threshold_uncertainty_score":0.0966863},"labels":[],"label_agreement":null},{"id":"W4327644549","doi":"10.1007/978-3-031-28238-6_24","title":"De-biasing Relevance Judgements for Fair Ranking","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); University of Waterloo; Toronto Metropolitan University","funders":"","keywords":"Ranking (information retrieval); Judgement; Relevance (law); Computer science; Set (abstract data type); Artificial intelligence; Biasing; Gender bias; Process (computing); Machine learning; Information retrieval; Psychology; Social psychology; Voltage","score_opus":0.03901985798360569,"score_gpt":0.2754367443815013,"score_spread":0.2364168863978956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644549","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059801103,0.0018003463,0.9843551,0.00091829355,0.00072253897,0.00021972905,0.00026605275,0.0010231577,0.004714704],"genre_scores_gemma":[0.29216465,0.0010952866,0.6894625,0.0010404055,0.0021192355,0.00075634755,0.0014450062,0.0010019834,0.010914625],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.934287,0.040353864,0.0038187413,0.0052496633,0.014563796,0.0017269183],"domain_scores_gemma":[0.84488225,0.107240096,0.0037911616,0.025688076,0.01699183,0.0014066692],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04927327,0.0018440393,0.0034448549,0.006000045,0.0025799335,0.00554489,0.004229721,0.0034935276,0.010842023],"category_scores_gemma":[0.20650418,0.0011594456,0.0021259845,0.005567659,0.0031854059,0.0076770047,0.0061154207,0.0060350066,0.0043882513],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011568352,0.00021036799,0.0017043202,0.0012430635,0.00053661177,0.00011535142,0.0006979773,0.027661666,0.007885227,0.13669881,0.035508394,0.7865813],"study_design_scores_gemma":[0.00020350193,0.0002163294,0.0011097427,0.00033785417,0.0002324422,0.00026355468,0.00014910373,0.32254365,0.00812242,0.6475799,0.01912905,0.000112442875],"about_ca_topic_score_codex":0.0017075771,"about_ca_topic_score_gemma":0.0035948406,"teacher_disagreement_score":0.04927327,"about_ca_system_score_codex":0.002308491,"about_ca_system_score_gemma":0.0031435322,"threshold_uncertainty_score":0.26058507},"labels":[],"label_agreement":null},{"id":"W4327644584","doi":"10.1007/978-3-031-28238-6_51","title":"Learning Query-Space Document Representations for High-Recall Retrieval","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; University of Waterloo; Toronto Metropolitan University","funders":"","keywords":"Computer science; Information retrieval; Set (abstract data type); Recall; Representation (politics); Space (punctuation); Range (aeronautics); Precision and recall; Artificial intelligence","score_opus":0.027236126660376232,"score_gpt":0.2840207212977938,"score_spread":0.2567845946374176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644584","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044018473,0.004138328,0.9396754,0.0006539817,0.00022722852,0.0001948488,0.0017616568,0.0074998112,0.0018302397],"genre_scores_gemma":[0.38116974,0.0034469687,0.59122765,0.00054550363,0.0006616867,0.0004982542,0.012288631,0.000628626,0.009532911],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9987062,0.00037803192,0.0001225151,0.00030621258,0.00031957738,0.00016752731],"domain_scores_gemma":[0.99776375,0.0011344912,0.00016712598,0.00046498055,0.00038377382,0.000085784035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001951683,0.0014191929,0.0019353457,0.0032867393,0.0005626069,0.0021085988,0.002096284,0.002144953,0.00539972],"category_scores_gemma":[0.006581736,0.0005767967,0.0014132169,0.003868878,0.0006431782,0.004024752,0.0015668336,0.0023506638,0.0042765187],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069368025,0.00038567977,0.0012335669,0.00032967018,0.00012919171,0.000092617796,0.00015022389,0.035933472,0.016910026,0.0067278533,0.023094865,0.9143191],"study_design_scores_gemma":[0.00009497979,0.00033254115,0.000836988,0.000061689345,0.00012474456,0.0002810393,0.0001310234,0.95982903,0.009719145,0.023313552,0.0052271485,0.000048124555],"about_ca_topic_score_codex":0.004597288,"about_ca_topic_score_gemma":0.005463167,"teacher_disagreement_score":0.00539972,"about_ca_system_score_codex":0.0012942472,"about_ca_system_score_gemma":0.0013036776,"threshold_uncertainty_score":0.018063903},"labels":[],"label_agreement":null},{"id":"W4327644606","doi":"10.1007/978-3-031-28238-6_57","title":"Neural Ad-Hoc Retrieval Meets Open Information Extraction","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto; University of Guelph; Toronto Metropolitan University","funders":"","keywords":"Computer science; Information retrieval; Information extraction; Task (project management); Range (aeronautics); Artificial intelligence; Relationship extraction; Artificial neural network; Natural language processing","score_opus":0.038498278859022225,"score_gpt":0.28689816379407646,"score_spread":0.24839988493505422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644606","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0108876685,0.006959418,0.9532976,0.0011176016,0.00072027807,0.00017819539,0.0007253795,0.004507144,0.02160678],"genre_scores_gemma":[0.21124451,0.008396041,0.70045394,0.0006022411,0.0018011493,0.00022649686,0.0041762777,0.0008881358,0.07221125],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988502,0.00021531543,0.00013905574,0.00026762957,0.00040057738,0.00012713864],"domain_scores_gemma":[0.9978625,0.00079377554,0.00009785302,0.00069250446,0.0004877901,0.00006544216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014157324,0.0009712569,0.001355578,0.0025389341,0.00089587213,0.0037758616,0.0017378771,0.0016448387,0.011210713],"category_scores_gemma":[0.0045216465,0.00066199293,0.0010416687,0.0036663185,0.00089186954,0.006828119,0.0026730823,0.0014930293,0.0121089555],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001409802,0.00012635457,0.00041843663,0.0004228037,0.00006503723,0.00010729536,0.00007712601,0.0064020893,0.018431112,0.026789676,0.026726032,0.92029303],"study_design_scores_gemma":[0.000046557714,0.00019241622,0.0012526175,0.00013294921,0.0002687012,0.0011866076,0.00022802602,0.5986288,0.07648691,0.20838687,0.1131109,0.00007866515],"about_ca_topic_score_codex":0.0017471786,"about_ca_topic_score_gemma":0.003262614,"teacher_disagreement_score":0.011210713,"about_ca_system_score_codex":0.0008847125,"about_ca_system_score_gemma":0.0011234665,"threshold_uncertainty_score":0.03750354},"labels":[],"label_agreement":null},{"id":"W4327644611","doi":"10.1007/978-3-031-28238-6_50","title":"Don’t Raise Your Voice, Improve Your Argument: Learning to Retrieve Convincing Arguments","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Toronto Metropolitan University","funders":"","keywords":"Computer science; Relevance (law); Information retrieval; Focus (optics); Content (measure theory); Argument (complex analysis); Value (mathematics); Artificial intelligence; Natural language processing; Machine learning","score_opus":0.02644597134160811,"score_gpt":0.2651516314128106,"score_spread":0.2387056600712025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327644611","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.069374055,0.0033960904,0.7849421,0.010284402,0.0016435462,0.00036319203,0.0007364349,0.0050100884,0.12425012],"genre_scores_gemma":[0.316048,0.0031667058,0.5931434,0.00305276,0.0010390227,0.00039389115,0.003509308,0.0022195813,0.07742735],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989236,0.0004483482,0.00005756518,0.00022292916,0.00029448097,0.00005303451],"domain_scores_gemma":[0.98210573,0.0153486505,0.0005512716,0.00082126644,0.00084667676,0.0003262966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020188668,0.0011726036,0.00061183254,0.000885022,0.0007658584,0.0041292454,0.0016630001,0.0023379126,0.029244578],"category_scores_gemma":[0.03275114,0.0005409564,0.0008001011,0.00063994667,0.0009732192,0.010928852,0.0021552849,0.004361908,0.022223579],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003478459,0.00031122993,0.0014114534,0.0006380697,0.000056250654,0.00029390544,0.0030032496,0.002055261,0.009358847,0.047475476,0.09007969,0.8449687],"study_design_scores_gemma":[0.00026797634,0.00048367694,0.0027356287,0.0011879144,0.00028080802,0.001956429,0.004556878,0.17133667,0.02517431,0.57691836,0.21495405,0.00014736086],"about_ca_topic_score_codex":0.00035878364,"about_ca_topic_score_gemma":0.00061290443,"teacher_disagreement_score":0.029244578,"about_ca_system_score_codex":0.0005572822,"about_ca_system_score_gemma":0.0008326807,"threshold_uncertainty_score":0.09783286},"labels":[],"label_agreement":null},{"id":"W4327645767","doi":"10.1007/978-3-031-28244-7_1","title":"Self-supervised Contrastive BERT Fine-tuning for Fusion-Based Reviewed-Item Retrieval","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Ranking (information retrieval); Artificial intelligence; Matching (statistics); Selection (genetic algorithm); Natural language processing; Information retrieval; Task (project management); Embedding; Machine learning","score_opus":0.03059425526710696,"score_gpt":0.25976128234527573,"score_spread":0.22916702707816877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327645767","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04501799,0.0024686302,0.94221735,0.00015858383,0.00023917007,0.00018045706,0.00043024882,0.005920604,0.003367017],"genre_scores_gemma":[0.5737909,0.0007173121,0.41100794,0.00043274555,0.00038229034,0.0002942898,0.0022723654,0.0006278949,0.010474196],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989477,0.0002058605,0.00007564349,0.00031149198,0.0003023011,0.00015702196],"domain_scores_gemma":[0.9986299,0.0005416503,0.000081371574,0.0002182577,0.00045874098,0.000070133035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017595214,0.0008296769,0.0018142988,0.0017290481,0.00055539946,0.00092513405,0.002309885,0.0012819098,0.0046714083],"category_scores_gemma":[0.0031593225,0.00037369409,0.000841341,0.001723092,0.00052635616,0.0015786017,0.0016349738,0.0011897456,0.0027508603],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008099151,0.00046094146,0.000992367,0.00021144447,0.00015541517,0.0000713173,0.0000822076,0.03149619,0.07000098,0.0016322598,0.009489745,0.8845973],"study_design_scores_gemma":[0.000044852626,0.00018961403,0.0015279999,0.000017599954,0.00010357581,0.0001401981,0.00003733134,0.97280824,0.019813986,0.002188648,0.003098437,0.00002949285],"about_ca_topic_score_codex":0.0038333612,"about_ca_topic_score_gemma":0.007551578,"teacher_disagreement_score":0.0046714083,"about_ca_system_score_codex":0.00058611284,"about_ca_system_score_gemma":0.0009269022,"threshold_uncertainty_score":0.015627384},"labels":[],"label_agreement":null},{"id":"W4327671617","doi":"10.48550/arxiv.2303.08518","title":"UPRISE: Universal Prompt Retrieval for Improving Zero-Shot Evaluation","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Task (project management); Computer science; Generalization; Set (abstract data type); Zero (linguistics); Shot (pellet); Labrador Retriever; Filling-in; Code (set theory); Computer security; Artificial intelligence; Medicine; Engineering; Programming language; Pathology; Mathematics","score_opus":0.20330816455176112,"score_gpt":0.24053851469883095,"score_spread":0.037230350147069824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4327671617","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06785131,0.0062184134,0.72832066,0.0008276002,0.0014741459,0.00095326087,0.0055903574,0.17696562,0.011798636],"genre_scores_gemma":[0.51000464,0.00096472923,0.44101247,0.0016761699,0.0005674648,0.0013127399,0.020013886,0.01178566,0.012662221],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993708,0.002879508,0.0004509786,0.0014062586,0.0011232562,0.0004319468],"domain_scores_gemma":[0.9904198,0.005580602,0.00027206205,0.0018981702,0.0013144192,0.00051493326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009021903,0.0035163707,0.0021166692,0.001990672,0.0008540001,0.0030709591,0.0035692928,0.002927127,0.013846084],"category_scores_gemma":[0.038348984,0.00068511616,0.0011107324,0.00093061797,0.0009320856,0.006322544,0.0054289876,0.003299338,0.008823624],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00295491,0.00068554,0.0041344794,0.0021626183,0.00052025815,0.00041982377,0.0009662327,0.033456013,0.040845986,0.006628782,0.12065249,0.7865729],"study_design_scores_gemma":[0.0007365589,0.0015580452,0.0036395504,0.00026779348,0.0002316777,0.0006361866,0.0005721376,0.87704146,0.05099869,0.031812787,0.032244425,0.0002606413],"about_ca_topic_score_codex":0.0037540975,"about_ca_topic_score_gemma":0.0068535395,"teacher_disagreement_score":0.013846084,"about_ca_system_score_codex":0.0012963816,"about_ca_system_score_gemma":0.0017048644,"threshold_uncertainty_score":0.04771298},"labels":[],"label_agreement":null},{"id":"W4328028649","doi":"10.1109/trustcom56396.2022.00098","title":"FLightNER: A Federated Learning Approach to Lightweight Named-Entity Recognition","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Federated learning; Premise; Named-entity recognition; Artificial intelligence; Machine learning; Task (project management)","score_opus":0.03284144225408836,"score_gpt":0.23148371673515603,"score_spread":0.19864227448106767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4328028649","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015030612,0.0004684516,0.93041897,0.0004639912,0.00015809006,0.00017240694,0.0018040114,0.0494929,0.0019904939],"genre_scores_gemma":[0.3434137,0.00040890553,0.6310025,0.0008547625,0.00014436104,0.0004020803,0.013020644,0.0019139829,0.008839101],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985661,0.00029523278,0.00011280723,0.0005902249,0.00031250287,0.00012317009],"domain_scores_gemma":[0.997182,0.0007113632,0.00013466415,0.0014087344,0.00042532285,0.00013789642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002998976,0.0012113947,0.0012908218,0.0013614349,0.00066098874,0.0018619843,0.003862606,0.0012972931,0.0037067563],"category_scores_gemma":[0.008035175,0.000596941,0.0015520864,0.0014008338,0.00062614883,0.0056121065,0.0028860276,0.0023192025,0.0030594103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007367951,0.0006514921,0.006839169,0.00025752114,0.00039132906,0.00033047958,0.0002540388,0.34453312,0.005688654,0.012500403,0.052683696,0.5751332],"study_design_scores_gemma":[0.00002196697,0.000046603887,0.0003254653,0.000012204489,0.000022236789,0.00007246319,0.000030788837,0.9789098,0.002939371,0.011698463,0.0059010955,0.000019497344],"about_ca_topic_score_codex":0.009190314,"about_ca_topic_score_gemma":0.013276549,"teacher_disagreement_score":0.009190314,"about_ca_system_score_codex":0.0010576759,"about_ca_system_score_gemma":0.0018012079,"threshold_uncertainty_score":0.018273652},"labels":[],"label_agreement":null},{"id":"W4352981330","doi":"10.1109/iscmi56532.2022.10068466","title":"Empirical Evaluation of Word Representation Methods in the Context of Candidate-Job Recommender Systems","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Durham College","funders":"","keywords":"Cosine similarity; Computer science; Ranking (information retrieval); Information retrieval; Similarity (geometry); tf–idf; Rank (graph theory); Recommender system; Context (archaeology); Matching (statistics); Set (abstract data type); Word (group theory); Artificial intelligence; Natural language processing; Term (time); Mathematics; Statistics; Pattern recognition (psychology)","score_opus":0.26316627833519207,"score_gpt":0.47511817729255207,"score_spread":0.21195189895736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4352981330","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8905608,0.0074093067,0.09280523,0.00056693214,0.00023402051,0.0005221277,0.0016999452,0.0019886014,0.0042131105],"genre_scores_gemma":[0.91003305,0.0007410021,0.08380988,0.000091631264,0.00010975564,0.0002480591,0.003389604,0.00013567082,0.0014413031],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98412573,0.010811814,0.0012260284,0.0015895445,0.0019030336,0.00034388228],"domain_scores_gemma":[0.9304078,0.05868451,0.0018908689,0.0037921155,0.004613937,0.00061067595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014762245,0.0011177092,0.0010212341,0.003728164,0.000924375,0.0019543504,0.0012294885,0.0019925286,0.0017703831],"category_scores_gemma":[0.067081496,0.00030544066,0.0007630127,0.003575052,0.0007223905,0.0034947512,0.0013996762,0.0010168091,0.0010856884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0046207956,0.0029552018,0.090534665,0.00313222,0.0013899155,0.000402605,0.0018634526,0.12160903,0.011938695,0.0028379678,0.0069592143,0.7517562],"study_design_scores_gemma":[0.000381534,0.004267857,0.051089086,0.00021150052,0.0005007547,0.00065934163,0.0020445941,0.91618747,0.015730822,0.0030285204,0.005727198,0.00017131689],"about_ca_topic_score_codex":0.006550811,"about_ca_topic_score_gemma":0.008038266,"teacher_disagreement_score":0.014762245,"about_ca_system_score_codex":0.0010550759,"about_ca_system_score_gemma":0.00083149335,"threshold_uncertainty_score":0.07807112},"labels":[],"label_agreement":null},{"id":"W4360764250","doi":"10.1109/icmla55696.2022.00030","title":"A Robust Approach to Fine-tune Pre-trained Transformer-based models for Text Summarization through Latent Space Compression","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Automatic summarization; Computer science; Encoder; Transformer; Autoencoder; Inference; Artificial intelligence; Deep learning; Speech recognition; Pattern recognition (psychology)","score_opus":0.07704970137916824,"score_gpt":0.2582683224406403,"score_spread":0.18121862106147207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4360764250","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004736115,0.00014428585,0.990529,0.000083987674,0.000046705543,0.0000395554,0.00014623605,0.0037769054,0.0004972801],"genre_scores_gemma":[0.18650286,0.00028067743,0.8033496,0.00026459675,0.00017259429,0.00025363817,0.002100829,0.0008904883,0.006184652],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99938,0.00016587762,0.000046982434,0.00019230417,0.00015474904,0.000060052145],"domain_scores_gemma":[0.9990157,0.00037423678,0.00007161737,0.00022740274,0.00027026873,0.00004089022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010924366,0.001490087,0.00081571453,0.00078866153,0.00033065895,0.00079124566,0.0016247341,0.0010329664,0.0035134596],"category_scores_gemma":[0.0028600825,0.00066486566,0.00110853,0.0007111239,0.000424176,0.0016882505,0.0010518272,0.0018510513,0.0029778243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033887368,0.00021905155,0.00067599246,0.0002022992,0.00022661492,0.0001844191,0.00016096706,0.29685616,0.10531802,0.00730524,0.008263227,0.5802491],"study_design_scores_gemma":[0.000018019426,0.00007165654,0.00020199796,0.000006153722,0.00003064454,0.000045218003,0.00001379604,0.96843755,0.026255865,0.0023148472,0.0025900074,0.000014172018],"about_ca_topic_score_codex":0.003481551,"about_ca_topic_score_gemma":0.0070504765,"teacher_disagreement_score":0.0035134596,"about_ca_system_score_codex":0.00059364137,"about_ca_system_score_gemma":0.0010151985,"threshold_uncertainty_score":0.011753678},"labels":[],"label_agreement":null},{"id":"W4360978668","doi":"10.1145/3581754.3584136","title":"Supporting Qualitative Analysis with Large Language Models: Combining Codebook with GPT-3 for Deductive Coding","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":207,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); Research Canada","funders":"","keywords":"Codebook; Generalizability theory; Computer science; Coding (social sciences); Task (project management); Natural language processing; Curiosity; Qualitative analysis; Qualitative research; Task analysis; Artificial intelligence; Qualitative property; Machine learning; Data science; Psychology; Social psychology","score_opus":0.08960805608801989,"score_gpt":0.39019869997724693,"score_spread":0.30059064388922707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4360978668","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003922875,0.00006566569,0.98720413,0.0006891242,0.00006374536,0.0009120545,0.0015476436,0.0041605174,0.0014342167],"genre_scores_gemma":[0.054076865,0.00005728307,0.93899685,0.00024670985,0.000023820887,0.002256502,0.0032068864,0.0005752765,0.0005597575],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.93974894,0.04490444,0.004260631,0.004840308,0.005573328,0.00067240506],"domain_scores_gemma":[0.67134106,0.24965946,0.008402717,0.04192705,0.026285108,0.0023846396],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.057022505,0.0020826003,0.0010921177,0.0047738855,0.0017906271,0.0067120567,0.005193062,0.0022157708,0.010095994],"category_scores_gemma":[0.32469282,0.0014585624,0.0022788576,0.0035998966,0.0036252972,0.010327976,0.012931254,0.005140884,0.006454557],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069183385,0.000399526,0.007821606,0.0034378602,0.00020638337,0.0006778925,0.057498004,0.045227405,0.011471382,0.16686064,0.03423222,0.6714752],"study_design_scores_gemma":[0.00028575608,0.0001656042,0.0010999725,0.0013037015,0.00006905326,0.00035412284,0.008000993,0.53758854,0.011600964,0.38863322,0.050668743,0.00022931998],"about_ca_topic_score_codex":0.007324775,"about_ca_topic_score_gemma":0.012135737,"teacher_disagreement_score":0.9429775,"about_ca_system_score_codex":0.005164192,"about_ca_system_score_gemma":0.009166841,"threshold_uncertainty_score":0.30156738},"labels":[],"label_agreement":null},{"id":"W4360989137","doi":"10.18280/ria.370130","title":"NeRBERT- A Biomedical Named Entity Recognition Tagger","year":2023,"lang":"fr","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Natural language processing; Computer science; Named-entity recognition; Artificial intelligence; Linguistics; Engineering; Philosophy","score_opus":0.12115317540638054,"score_gpt":0.30755804131551046,"score_spread":0.1864048659091299,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4360989137","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.072906524,0.0036714307,0.7052241,0.0015973173,0.0011211309,0.0006416727,0.0393677,0.16143574,0.014034427],"genre_scores_gemma":[0.2482681,0.0018034395,0.5952268,0.0011780264,0.00018814829,0.0006886985,0.11009728,0.0024822503,0.0400673],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948704,0.00010449967,0.000040589137,0.00020031749,0.00011795942,0.000049491024],"domain_scores_gemma":[0.9990491,0.0003786072,0.00006556485,0.00023681192,0.00022791122,0.000041925345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012644057,0.0011386268,0.00072709634,0.0017892749,0.00051572017,0.0010134692,0.0016183362,0.0014500755,0.0043742093],"category_scores_gemma":[0.002665082,0.0004127283,0.00095427106,0.0013799446,0.0003019454,0.0032262234,0.0011984406,0.0012575974,0.006775367],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009124762,0.0003688995,0.006253205,0.0010012783,0.0004511574,0.00076944375,0.0002889217,0.07260499,0.030655349,0.010035195,0.14915892,0.7275003],"study_design_scores_gemma":[0.00012146986,0.0004376559,0.005133782,0.00014479177,0.0002671586,0.0013612881,0.00023573026,0.7666988,0.085408665,0.014724587,0.12528743,0.00017867499],"about_ca_topic_score_codex":0.006994069,"about_ca_topic_score_gemma":0.014623934,"teacher_disagreement_score":0.006994069,"about_ca_system_score_codex":0.0009677971,"about_ca_system_score_gemma":0.0013477351,"threshold_uncertainty_score":0.014633179},"labels":[],"label_agreement":null},{"id":"W4361281614","doi":"10.31219/osf.io/wuzy9","title":"Towards Automated Assessment of Scientific Explanations in Turkish using Language Transfer","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation","keywords":"Turkish; Computer science; Natural language processing; Formative assessment; Annotation; Artificial intelligence; Hebrew; Transformer; Language model; Transfer of learning; Linguistics; Mathematics education; Engineering; Psychology","score_opus":0.09673746239018131,"score_gpt":0.3736425509112112,"score_spread":0.27690508852102985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4361281614","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16533343,0.0005038179,0.81144774,0.00092339056,0.000084936604,0.0006132285,0.0014498024,0.014193771,0.0054499423],"genre_scores_gemma":[0.65912515,0.00021834477,0.33412296,0.00011130614,0.00004028201,0.00033056282,0.0032219402,0.00035304332,0.0024763849],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99471176,0.0034987128,0.00023052999,0.00083508313,0.00057315273,0.00015073866],"domain_scores_gemma":[0.98212147,0.011541924,0.0013203108,0.0015316455,0.0031631403,0.00032150748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051188553,0.0012816491,0.0004106594,0.0032776948,0.00056313805,0.0026563685,0.0011652181,0.000995267,0.0030867322],"category_scores_gemma":[0.02359469,0.00029654853,0.0010309966,0.0012178272,0.0006372341,0.0037516707,0.0023864475,0.0018785184,0.0022539212],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050979864,0.00045002362,0.015952561,0.0008050044,0.00016170209,0.00046294797,0.007947007,0.028047487,0.033559844,0.008695789,0.008011692,0.895396],"study_design_scores_gemma":[0.00009185879,0.00054699986,0.018432105,0.00023551445,0.00020332454,0.000622252,0.0073156646,0.8462713,0.070371,0.028849635,0.026869934,0.00019046584],"about_ca_topic_score_codex":0.0032897368,"about_ca_topic_score_gemma":0.0041205683,"teacher_disagreement_score":0.0051188553,"about_ca_system_score_codex":0.0014683994,"about_ca_system_score_gemma":0.0018154926,"threshold_uncertainty_score":0.027071476},"labels":[],"label_agreement":null},{"id":"W4361856293","doi":"10.1037/xlm0001226","title":"Modeling verbal short-term memory: A walk around the neighborhood.","year":2023,"lang":"en","type":"article","venue":"Journal of Experimental Psychology Learning Memory and Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Optimal distinctiveness theory; Recall; Psychology; Term (time); Cognitive psychology; Short-term memory; Semantic memory; PsycINFO; Episodic memory; Semantics (computer science); Cognition; Computer science; Social psychology; Working memory","score_opus":0.05058692894680105,"score_gpt":0.33380268679912745,"score_spread":0.2832157578523264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4361856293","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19664976,0.0047862823,0.7818957,0.002253752,0.00012199584,0.00014604292,0.00047967624,0.00028696258,0.013379795],"genre_scores_gemma":[0.86259395,0.0025670445,0.12936138,0.00019012977,0.00014804554,0.00035548027,0.00041230573,0.00011076874,0.0042608944],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995888,0.00022903208,0.000019693205,0.00009395461,0.000043593143,0.000025004758],"domain_scores_gemma":[0.9973092,0.001781075,0.00027352665,0.0003225905,0.00016327432,0.0001503574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014531526,0.00071674684,0.0008641256,0.0011376614,0.00058297336,0.0020839612,0.0026877937,0.0013436243,0.0030182377],"category_scores_gemma":[0.007705449,0.00053707184,0.0013067344,0.0008357108,0.0010502912,0.00446862,0.0017645909,0.0014807112,0.0006166647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019875051,0.00018126487,0.014840822,0.00032262778,0.00029598986,0.0007674,0.001837764,0.5219326,0.0027424365,0.41492894,0.0017521356,0.040199243],"study_design_scores_gemma":[0.00001680182,0.00009014021,0.0010539715,0.000036171616,0.000051793686,0.00013359672,0.000112822796,0.75074965,0.00023492615,0.24540104,0.0020999142,0.000019267716],"about_ca_topic_score_codex":0.006124572,"about_ca_topic_score_gemma":0.0047347373,"teacher_disagreement_score":0.006124572,"about_ca_system_score_codex":0.00071467715,"about_ca_system_score_gemma":0.00056970556,"threshold_uncertainty_score":0.012177885},"labels":[],"label_agreement":null},{"id":"W4362559515","doi":"10.3390/electronics12071692","title":"SS-BERT: A Semantic Information Selecting Approach for Open-Domain Question Answering","year":2023,"lang":"en","type":"article","venue":"Electronics","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"National Natural Science Foundation of China","keywords":"Question answering; Computer science; Open domain; Information retrieval; Labrador Retriever; Domain (mathematical analysis); Precision and recall; Selection (genetic algorithm); Encoder; Dual (grammatical number); Artificial intelligence; Natural language processing; Mathematics; Medicine","score_opus":0.019301273658451988,"score_gpt":0.27453031527493327,"score_spread":0.2552290416164813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362559515","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016987389,0.0022357637,0.9553421,0.0008363626,0.00020466775,0.0005585144,0.0021347543,0.01615883,0.005541622],"genre_scores_gemma":[0.3798304,0.0014314075,0.5917743,0.0009651212,0.00046547267,0.0006644673,0.010119609,0.0007905537,0.0139587615],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99824893,0.0007340292,0.00010002442,0.00043762097,0.0003898846,0.000089461784],"domain_scores_gemma":[0.99754155,0.0013602194,0.00011104736,0.0003406643,0.00049639714,0.00015020353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025954675,0.0015937074,0.0010002995,0.0038132658,0.00067638635,0.0012260058,0.0020752796,0.0018437075,0.0063703367],"category_scores_gemma":[0.0063278005,0.0005101395,0.0014598353,0.001985405,0.0009258699,0.0032907468,0.001694917,0.0019999652,0.0035621012],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072259817,0.0006772149,0.0039212364,0.0012779924,0.00025221685,0.00042137422,0.0008648863,0.06106513,0.023683177,0.030183094,0.06878456,0.8081465],"study_design_scores_gemma":[0.00010220205,0.00030710877,0.0013719157,0.00007186698,0.000121512596,0.00035761195,0.0001570528,0.91946304,0.01048691,0.034153152,0.033339452,0.0000680767],"about_ca_topic_score_codex":0.0065298798,"about_ca_topic_score_gemma":0.012944845,"teacher_disagreement_score":0.0065298798,"about_ca_system_score_codex":0.0013145581,"about_ca_system_score_gemma":0.0016474309,"threshold_uncertainty_score":0.021310925},"labels":[],"label_agreement":null},{"id":"W4362575620","doi":"10.22215/etd/2023-15426","title":"Knowledge Graph Generation for Unstructured Data Using Data Processing Pipeline","year":2023,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Unstructured data; Computer science; Pipeline (software); Information extraction; Knowledge graph; Coreference; Graph; Information retrieval; The Internet; Knowledge extraction; Data mining; Natural language processing; Artificial intelligence; Data science; Resolution (logic); Big data; World Wide Web; Theoretical computer science; Programming language","score_opus":0.269369695870474,"score_gpt":0.4050623587275948,"score_spread":0.13569266285712078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362575620","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017334709,0.00013812917,0.945179,0.00045716017,0.00005376096,0.0004312424,0.0042450847,0.029979512,0.0021813982],"genre_scores_gemma":[0.10436117,0.00015316674,0.8743642,0.00014253093,0.000022533053,0.00038516274,0.01778226,0.000856329,0.0019326368],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915373,0.00014162317,0.000071052375,0.0003161983,0.00025913378,0.000058164613],"domain_scores_gemma":[0.99702877,0.0014936592,0.00011692897,0.00072224304,0.00055854756,0.0000799526],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014821197,0.0009816567,0.00062732556,0.0037450753,0.0007553653,0.0014841781,0.0015824159,0.0008912117,0.004358451],"category_scores_gemma":[0.006270119,0.00054140104,0.0016959155,0.0032182385,0.0005112289,0.003081447,0.0021340898,0.0015230306,0.0020446642],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032836956,0.0004447259,0.004764374,0.0006912994,0.00018198269,0.0005591401,0.0009318975,0.06968354,0.021706836,0.027399814,0.045345496,0.8279625],"study_design_scores_gemma":[0.00005926092,0.000084757114,0.001302905,0.00004458062,0.00006895155,0.00018943827,0.00029683605,0.8902264,0.027796116,0.049906522,0.029981984,0.000042213556],"about_ca_topic_score_codex":0.010679491,"about_ca_topic_score_gemma":0.017405251,"teacher_disagreement_score":0.010679491,"about_ca_system_score_codex":0.0011117868,"about_ca_system_score_gemma":0.0019599716,"threshold_uncertainty_score":0.021234632},"labels":[],"label_agreement":null},{"id":"W4362598841","doi":"10.48550/arxiv.2304.01019","title":"Simple Yet Effective Neural Ranking and Reranking Baselines for Cross-Lingual Information Retrieval","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Ranking (information retrieval); Information retrieval; Pace; Context (archaeology); Task (project management); Simple (philosophy); Natural language processing; Artificial intelligence; Query expansion","score_opus":0.07533395319001727,"score_gpt":0.2354221850205383,"score_spread":0.160088231830521,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362598841","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13751017,0.014549075,0.7916204,0.0015392003,0.0015255983,0.0009409464,0.004798604,0.029500213,0.0180158],"genre_scores_gemma":[0.45993036,0.0020351368,0.5125501,0.00040891903,0.00054911216,0.00080220896,0.012936834,0.0013091682,0.009478154],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99414116,0.0024437073,0.0004953362,0.001317531,0.0012115911,0.00039068083],"domain_scores_gemma":[0.99278957,0.0023812442,0.0002934737,0.0020613207,0.0022924554,0.0001819216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010128313,0.0022256032,0.0016746934,0.003921855,0.0016551685,0.0025267343,0.00365482,0.0022312298,0.00434698],"category_scores_gemma":[0.021039665,0.000772196,0.0011790172,0.003606902,0.00078614,0.005970321,0.0020290646,0.0033774117,0.004310065],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000858873,0.0009610637,0.0034925449,0.0006435693,0.0005239364,0.00008664501,0.0002117977,0.11397776,0.01436313,0.008114314,0.033053793,0.82371265],"study_design_scores_gemma":[0.0001699462,0.0004782222,0.0029529885,0.00008553397,0.00018104589,0.00009733396,0.0001117411,0.9567474,0.015670652,0.01453604,0.008869955,0.000099078636],"about_ca_topic_score_codex":0.015574264,"about_ca_topic_score_gemma":0.032730684,"teacher_disagreement_score":0.015574264,"about_ca_system_score_codex":0.0019191077,"about_ca_system_score_gemma":0.0022097195,"threshold_uncertainty_score":0.05356431},"labels":[],"label_agreement":null},{"id":"W4362601647","doi":"10.1001/jamasurg.2023.0958","title":"Errors in Author Names and Affiliation","year":2023,"lang":"en","type":"erratum","venue":"JAMA Surgery","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Medicine; Gerontology","score_opus":0.04197408462975459,"score_gpt":0.2712112143318795,"score_spread":0.2292371297021249,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362601647","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008976796,0.0027368623,0.011628121,0.10839071,0.68397397,0.0005212175,0.09343174,0.004807605,0.09361214],"genre_scores_gemma":[0.028833596,0.008403448,0.029201748,0.08043119,0.05341994,0.0031846,0.08189563,0.009806901,0.7048231],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97432804,0.005577253,0.005916015,0.0033072094,0.009328876,0.0015426253],"domain_scores_gemma":[0.86177975,0.04248639,0.0117828725,0.018342303,0.06258139,0.003027372],"candidate_categories":["research_integrity","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.014511316,0.0018694852,0.0026572784,0.005954605,0.004934526,0.0054540113,0.0031369836,0.004116148,0.32529786],"category_scores_gemma":[0.21855505,0.0015016476,0.0013954266,0.015909653,0.001984297,0.0045002806,0.003890668,0.009345285,0.34107882],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035597554,0.000004617822,0.00008805934,0.00008270082,0.0000053317635,0.000027845206,0.00008468422,0.000053649503,0.000027684351,0.0024996207,0.98997223,0.007117979],"study_design_scores_gemma":[0.000034966968,0.000011809577,0.00030204433,0.000423082,0.0000129001155,0.000095422365,0.00020885855,0.0001458873,0.00025439728,0.0037058417,0.9947817,0.00002311129],"about_ca_topic_score_codex":0.010522496,"about_ca_topic_score_gemma":0.011281812,"teacher_disagreement_score":0.9958838,"about_ca_system_score_codex":0.004107957,"about_ca_system_score_gemma":0.012349394,"threshold_uncertainty_score":0.9623807},"labels":[],"label_agreement":null},{"id":"W4362654035","doi":"10.1007/978-3-031-29168-5_4","title":"COLIEE 2022 Summary: Methods for Legal Document Retrieval and Entailment","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Task (project management); Computer science; Statute; Component (thermodynamics); Logical consequence; Textual entailment; Competition (biology); Information retrieval; Natural language processing; Artificial intelligence; Law; Political science","score_opus":0.02988480079389599,"score_gpt":0.32232969507133435,"score_spread":0.29244489427743836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362654035","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016410605,0.03547264,0.66908324,0.022055896,0.024417276,0.0008314154,0.023144657,0.024201088,0.19915275],"genre_scores_gemma":[0.01675972,0.012969578,0.30007935,0.005007065,0.013883096,0.00077396334,0.05509336,0.012163871,0.58326995],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975453,0.0003795037,0.00015211762,0.00038216388,0.0014061425,0.00013477336],"domain_scores_gemma":[0.99404645,0.0020845376,0.00008437519,0.0007267121,0.002812324,0.0002456108],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032970933,0.0018313995,0.0012721102,0.0046819537,0.0014433352,0.004671743,0.0028066786,0.0021859154,0.11281858],"category_scores_gemma":[0.012248622,0.00076536776,0.001257719,0.004963053,0.00087584316,0.0046108896,0.0013032592,0.0028819484,0.06259366],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006418715,0.00002275171,0.000056611858,0.00033227704,0.000020235751,0.000019565112,0.00003537275,0.000636453,0.00073451165,0.008901558,0.7923261,0.19685048],"study_design_scores_gemma":[0.000039694056,0.000035365854,0.0004536356,0.00014631294,0.000043541168,0.00011595928,0.00003153701,0.005491049,0.004816382,0.020880537,0.9678955,0.00005046877],"about_ca_topic_score_codex":0.025154853,"about_ca_topic_score_gemma":0.032748304,"teacher_disagreement_score":0.11281858,"about_ca_system_score_codex":0.0028877023,"about_ca_system_score_gemma":0.002793007,"threshold_uncertainty_score":0.37741572},"labels":[],"label_agreement":null},{"id":"W4362700380","doi":"10.1007/s11227-023-05209-z","title":"Exploring implicit persona knowledge for personalized dialogue generation","year":2023,"lang":"en","type":"article","venue":"The Journal of Supercomputing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Social Science Fund of China","keywords":"Persona; Computer science; Consistency (knowledge bases); Key (lock); Human–computer interaction; Style (visual arts); World Wide Web; Data science; Artificial intelligence","score_opus":0.25661851092955423,"score_gpt":0.3163193920778916,"score_spread":0.05970088114833738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4362700380","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07905271,0.00084981846,0.90583193,0.00099317,0.00021965316,0.00022510948,0.0009672817,0.0039009696,0.007959268],"genre_scores_gemma":[0.824898,0.00033112484,0.16691218,0.00021747821,0.00013087533,0.00022565576,0.0020806442,0.0004562212,0.0047477856],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973341,0.0014783951,0.0000992689,0.0006113278,0.0003105088,0.00016638798],"domain_scores_gemma":[0.9912747,0.0070919446,0.00018867392,0.00063962716,0.00058783044,0.00021723483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026396748,0.00093317777,0.00075752387,0.0012040088,0.0009692707,0.0025061045,0.0015667097,0.0019312946,0.007351538],"category_scores_gemma":[0.016391618,0.0007239473,0.0009968674,0.000847983,0.0006129392,0.004629547,0.0027968786,0.0025041967,0.0025235154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023700288,0.0009898508,0.011988635,0.0010856735,0.00044533773,0.0009255678,0.009680212,0.07343304,0.030400002,0.038729623,0.02017899,0.809773],"study_design_scores_gemma":[0.000085609936,0.0001777242,0.0020997995,0.00011437417,0.00015213464,0.00030828468,0.0012616946,0.92565817,0.009331107,0.04984784,0.010890665,0.00007254494],"about_ca_topic_score_codex":0.003028693,"about_ca_topic_score_gemma":0.0037721756,"teacher_disagreement_score":0.007351538,"about_ca_system_score_codex":0.0006743242,"about_ca_system_score_gemma":0.0009519826,"threshold_uncertainty_score":0.024593353},"labels":[],"label_agreement":null},{"id":"W4366127141","doi":"10.1007/978-3-031-25759-9_12","title":"Human-Centric Question-Answering System with Linguistic Terms","year":2023,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Question answering; Computer science; Artificial intelligence; Natural language; Natural (archaeology); Natural language processing","score_opus":0.10774633212867449,"score_gpt":0.3556728525025759,"score_spread":0.24792652037390142,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366127141","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05266939,0.0013818073,0.8777651,0.00084580114,0.00020967635,0.0004605664,0.0030788833,0.046988197,0.016600652],"genre_scores_gemma":[0.33066598,0.00052681816,0.6324762,0.00082364463,0.00023882832,0.00043920297,0.007926595,0.0008518805,0.026050854],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993781,0.00015231225,0.00004504472,0.00024558927,0.00013091913,0.00004807114],"domain_scores_gemma":[0.9993698,0.00027194768,0.00002690788,0.00013603523,0.00013385408,0.00006142867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010756478,0.00056290755,0.0011059006,0.0011263861,0.00071511895,0.0016547581,0.0016061127,0.0012271613,0.010161285],"category_scores_gemma":[0.001611095,0.00034123324,0.0006902763,0.001297148,0.00035694425,0.00323169,0.0015895133,0.0008417973,0.006195573],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013225662,0.0007959935,0.0035083555,0.0007754061,0.00025337946,0.00058453024,0.0013479224,0.007621367,0.15360026,0.027546493,0.10070711,0.7019366],"study_design_scores_gemma":[0.0002980081,0.0004603867,0.004806629,0.00008177696,0.0005256503,0.0012722315,0.00066029665,0.6938575,0.12596457,0.053900376,0.11804081,0.00013174485],"about_ca_topic_score_codex":0.0021248055,"about_ca_topic_score_gemma":0.0025674955,"teacher_disagreement_score":0.010161285,"about_ca_system_score_codex":0.0006344407,"about_ca_system_score_gemma":0.0010138752,"threshold_uncertainty_score":0.033992887},"labels":[],"label_agreement":null},{"id":"W4366201187","doi":"10.1109/access.2023.3267746","title":"B-NER: A Novel Bangla Named Entity Recognition Dataset With Largest Entities and Its Baseline Evaluation","year":2023,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Named-entity recognition; Computer science; Bengali; Natural language processing; Artificial intelligence; Entity linking; Baseline (sea); Benchmark (surveying); Sentence; F1 score; Task (project management); Information retrieval","score_opus":0.13697073902502488,"score_gpt":0.34404632604811985,"score_spread":0.20707558702309498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366201187","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22748798,0.0060082404,0.06450721,0.0021100978,0.002048588,0.0029422792,0.62398595,0.033431157,0.037478495],"genre_scores_gemma":[0.06205065,0.0006745115,0.05637589,0.00050762517,0.00010270786,0.0013308057,0.87193984,0.0005598449,0.0064581567],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9963408,0.0006273408,0.0006908786,0.0012858713,0.0008168415,0.00023823096],"domain_scores_gemma":[0.99591357,0.000832755,0.000298431,0.0015458341,0.0011830649,0.0002262493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002922909,0.0019454438,0.0014126074,0.0036426953,0.002263663,0.0019185297,0.0033208558,0.0025215081,0.005100152],"category_scores_gemma":[0.006926817,0.00038075924,0.0016860146,0.0034139191,0.0010043738,0.0038245309,0.0024865167,0.0016347832,0.0073574665],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002388242,0.0026018089,0.023119975,0.007465513,0.00072240294,0.0022936647,0.0010817712,0.025235016,0.045164615,0.008774351,0.5869299,0.2942227],"study_design_scores_gemma":[0.00055724336,0.0010627004,0.07530535,0.0010149787,0.0005292509,0.005627116,0.0026301749,0.118892275,0.062750466,0.007030364,0.7240108,0.0005893386],"about_ca_topic_score_codex":0.013639491,"about_ca_topic_score_gemma":0.02279222,"teacher_disagreement_score":0.013639491,"about_ca_system_score_codex":0.0015975,"about_ca_system_score_gemma":0.0015999456,"threshold_uncertainty_score":0.027120173},"labels":[],"label_agreement":null},{"id":"W4366549767","doi":"10.1145/3544548.3580895","title":"Enabling Conversational Interaction with Mobile UI using Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":141,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Human–computer interaction; Task (project management); Mobile device; Natural language; Language model; Multimedia; Artificial intelligence; World Wide Web","score_opus":0.04095757817540156,"score_gpt":0.29244990942456944,"score_spread":0.25149233124916787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366549767","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029254336,0.0004543133,0.9410185,0.00032053486,0.00007305972,0.0002601173,0.0006997134,0.02626429,0.0016551352],"genre_scores_gemma":[0.41907972,0.00042068688,0.569923,0.00041447216,0.00011117685,0.0009195388,0.00379317,0.0016614547,0.003676817],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9980447,0.000998221,0.00011675513,0.00048619704,0.0002574225,0.00009661458],"domain_scores_gemma":[0.9945201,0.0040794373,0.00017389105,0.00070401636,0.0003292476,0.00019326028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025678547,0.0019417298,0.0009913038,0.00071260356,0.0006298456,0.0018058177,0.001964612,0.0013493096,0.0032234907],"category_scores_gemma":[0.011596379,0.0007534483,0.0017797592,0.00048489068,0.0005498564,0.0034958806,0.0031470337,0.002317374,0.002754456],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001286743,0.00078141055,0.0062221964,0.001303514,0.00046901553,0.0010899262,0.005370328,0.18512133,0.09889018,0.01009422,0.025593884,0.6637773],"study_design_scores_gemma":[0.000056030214,0.00016310475,0.00064964494,0.0000453413,0.000056554065,0.00018063244,0.00042371542,0.9604011,0.015203077,0.0100683775,0.012679875,0.00007253041],"about_ca_topic_score_codex":0.0039141728,"about_ca_topic_score_gemma":0.006580844,"teacher_disagreement_score":0.0039141728,"about_ca_system_score_codex":0.000699098,"about_ca_system_score_gemma":0.0009890792,"threshold_uncertainty_score":0.013580263},"labels":[],"label_agreement":null},{"id":"W4366966678","doi":"10.1109/wi-iat55865.2022.00024","title":"NLI-based Filtering for Data Augmentation in Topic Classification","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wilfrid Laurier University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Leverage (statistics); Artificial intelligence; Machine learning; Task (project management); Filter (signal processing); Inference; Generalization; Encoder; Data mining; Language model; Natural language processing","score_opus":0.2319445510874427,"score_gpt":0.34777021792235907,"score_spread":0.11582566683491638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366966678","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007990923,0.0005124101,0.9857255,0.00033936749,0.00010646854,0.00026297115,0.0004713662,0.003848403,0.0007425762],"genre_scores_gemma":[0.13125066,0.00035606525,0.8602511,0.00056184904,0.00025602896,0.0009799246,0.0041599856,0.0006096368,0.0015748242],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948355,0.0023082464,0.0003859095,0.0014740642,0.00073707954,0.0002592493],"domain_scores_gemma":[0.97949445,0.01285564,0.0008441132,0.004807019,0.0017023706,0.00029649277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011946709,0.0016487585,0.0018249578,0.0028727644,0.001596245,0.0020863346,0.00335963,0.0021514394,0.0037238894],"category_scores_gemma":[0.033300072,0.00087055645,0.0025952095,0.002847891,0.0017114158,0.004489403,0.0028605626,0.0040169074,0.003259085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076348934,0.00072306505,0.010110637,0.00078796933,0.00025601848,0.00019015178,0.0012685008,0.05610155,0.023494555,0.020866986,0.016901985,0.8685351],"study_design_scores_gemma":[0.000118365126,0.00022945263,0.003104805,0.00013600364,0.00011643075,0.00027995568,0.00023643594,0.9187123,0.021631708,0.037381463,0.017976318,0.000076744],"about_ca_topic_score_codex":0.004078442,"about_ca_topic_score_gemma":0.0095157195,"teacher_disagreement_score":0.011946709,"about_ca_system_score_codex":0.0013966451,"about_ca_system_score_gemma":0.002684503,"threshold_uncertainty_score":0.06318098},"labels":[],"label_agreement":null},{"id":"W4366966846","doi":"10.1109/wi-iat55865.2022.00118","title":"A Noise-enhanced Fuse Model for Passage Ranking","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Ranking (information retrieval); Fuse (electrical); Artificial intelligence; Deep learning; Task (project management); Noise (video); Language model; Field (mathematics); Artificial neural network; Machine learning; Natural language processing; Deep neural networks; Engineering","score_opus":0.03584979530407897,"score_gpt":0.25647189010875093,"score_spread":0.22062209480467196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366966846","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0873604,0.0045238263,0.8905519,0.0010547144,0.0005659801,0.00014720697,0.0016162058,0.0066407225,0.0075389813],"genre_scores_gemma":[0.8343945,0.0016769122,0.13395365,0.0005508011,0.0006167803,0.00022198926,0.005314513,0.00044622767,0.02282461],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994555,0.00010340215,0.000029991179,0.00019238402,0.00012482307,0.000093873474],"domain_scores_gemma":[0.99939585,0.00020190388,0.000056865592,0.00007179701,0.00022065235,0.00005292435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013089591,0.0014157136,0.0015742473,0.0018283972,0.0007138165,0.0014955741,0.0017538538,0.0016058343,0.0029412322],"category_scores_gemma":[0.0026705703,0.0004180534,0.0013994111,0.001698766,0.0006415338,0.002536292,0.0008991131,0.0015266568,0.002249957],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064066256,0.000249593,0.0037222637,0.0001838863,0.00023995983,0.00029579288,0.00022953008,0.6840953,0.010933714,0.016626911,0.016502839,0.26627943],"study_design_scores_gemma":[0.000010465729,0.000039834944,0.00023364562,0.000004772766,0.000025482766,0.000032039334,0.0000069661164,0.995372,0.00077622355,0.0024236522,0.0010638259,0.00001107444],"about_ca_topic_score_codex":0.020283496,"about_ca_topic_score_gemma":0.020308258,"teacher_disagreement_score":0.020283496,"about_ca_system_score_codex":0.0012966378,"about_ca_system_score_gemma":0.0011558568,"threshold_uncertainty_score":0.040330827},"labels":[],"label_agreement":null},{"id":"W4366967218","doi":"10.1109/wi-iat55865.2022.00141","title":"A Multi-Dimensional Semantic Pseudo-Relevance Feedback Information Retrieval Model","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"China Scholarship Council; National Natural Science Foundation of China","keywords":"Computer science; Information retrieval; Ranking (information retrieval); Paragraph; Relevance (law); Relevance feedback; Sentence; Semantic similarity; Polysemy; Concept search; Word (group theory); Artificial intelligence; Natural language processing; Image retrieval; Search engine; Web search query; World Wide Web","score_opus":0.023053501409830972,"score_gpt":0.2412059020229618,"score_spread":0.21815240061313082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366967218","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030624075,0.0017288107,0.9587205,0.0011268505,0.00015292237,0.0002263249,0.00040252405,0.0011781395,0.005839847],"genre_scores_gemma":[0.82484245,0.0015370015,0.1565038,0.00046017897,0.00035753124,0.0005704606,0.0007765184,0.000100552934,0.014851573],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99753654,0.00081779604,0.00016853515,0.00048646092,0.00081239047,0.00017824741],"domain_scores_gemma":[0.9985489,0.0005324232,0.00014588848,0.00013213554,0.0005832404,0.000057325422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024789,0.0010365389,0.0015866531,0.0018859914,0.0006719173,0.0014137358,0.0027580238,0.0016010976,0.0036962754],"category_scores_gemma":[0.004466234,0.0004985943,0.001104675,0.0023677049,0.00080921245,0.004020213,0.0009385965,0.0010492117,0.0014995629],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073731627,0.00049966737,0.0029855757,0.00079803,0.00026525874,0.00061894336,0.0007229884,0.54028094,0.013982154,0.12036385,0.013686788,0.30505842],"study_design_scores_gemma":[0.000042721636,0.00009254886,0.0002577978,0.000008356186,0.00003398801,0.00012101606,0.0000146497405,0.9801855,0.00064407184,0.01719216,0.0013833841,0.000023844072],"about_ca_topic_score_codex":0.0052861054,"about_ca_topic_score_gemma":0.003839769,"teacher_disagreement_score":0.0052861054,"about_ca_system_score_codex":0.0014527765,"about_ca_system_score_gemma":0.0011881307,"threshold_uncertainty_score":0.013109803},"labels":[],"label_agreement":null},{"id":"W4366967907","doi":"10.1109/wi-iat55865.2022.00046","title":"Entity Level QA Pairs Dataset for Sentiment Analysis","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Sentiment analysis; Task (project management); Character (mathematics); Series (stratigraphy); Information retrieval; Span (engineering); Natural language processing; Artificial intelligence; Gauge (firearms); Mathematics","score_opus":0.08130547624460353,"score_gpt":0.2964838357407647,"score_spread":0.21517835949616115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366967907","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019568617,0.0007308652,0.004303205,0.000566672,0.00025029868,0.0005517965,0.962138,0.0048327404,0.007057885],"genre_scores_gemma":[0.017857991,0.00014744893,0.00783634,0.00014551917,0.00006125951,0.00042828702,0.9713212,0.00011480158,0.0020872],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987539,0.00028672616,0.0001947191,0.00025759707,0.00037087436,0.0001362324],"domain_scores_gemma":[0.9984724,0.00031426008,0.00017839184,0.00031272936,0.0005740202,0.00014812704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010117639,0.001662208,0.0007009426,0.0034563406,0.00090636796,0.0010068503,0.0011850682,0.0012360029,0.011260732],"category_scores_gemma":[0.0035803467,0.00024523336,0.0011088849,0.0026486986,0.00024728756,0.0016468388,0.0014170712,0.0010939831,0.013896066],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005538007,0.00038034626,0.0076569826,0.0015509493,0.00014499399,0.00035277568,0.0003453923,0.0020223882,0.009712765,0.0035863847,0.92772806,0.045965277],"study_design_scores_gemma":[0.00030942468,0.00031810903,0.037246387,0.0002462727,0.00009567675,0.00067404076,0.0009850138,0.03015158,0.009854728,0.0053521465,0.91464996,0.00011661879],"about_ca_topic_score_codex":0.008467915,"about_ca_topic_score_gemma":0.0148191815,"teacher_disagreement_score":0.011260732,"about_ca_system_score_codex":0.00097021536,"about_ca_system_score_gemma":0.0008545085,"threshold_uncertainty_score":0.03767085},"labels":[],"label_agreement":null},{"id":"W4367047001","doi":"10.1145/3543507.3583265","title":"Learning Denoised and Interpretable Session Representation for Conversational Search","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Interpretability; Computer science; Session (web analytics); Artificial intelligence; Representation (politics); Information retrieval; Machine learning; Search engine; Encoder; Natural language processing; World Wide Web","score_opus":0.052936613035745636,"score_gpt":0.33052167384140346,"score_spread":0.2775850608056578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4367047001","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044840116,0.0011562686,0.94662017,0.00026094328,0.000072883675,0.00012352022,0.00080212625,0.0041437494,0.001980246],"genre_scores_gemma":[0.64762175,0.0009175007,0.3334323,0.0005586402,0.00019428768,0.0003961436,0.0056842547,0.0006244451,0.010570635],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904066,0.00026721356,0.000049369948,0.0003394464,0.00019148558,0.000111755624],"domain_scores_gemma":[0.9988695,0.00049167237,0.000073674346,0.00029677738,0.00020104722,0.00006737129],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009353978,0.00095219444,0.0010144105,0.0010246973,0.0004269925,0.000880665,0.0017447959,0.0011505592,0.0024577382],"category_scores_gemma":[0.0046570813,0.0003719213,0.001048751,0.00091559463,0.00060916296,0.0025364235,0.00157065,0.0016071154,0.0020943694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001265474,0.0005150858,0.0038305698,0.00064928964,0.00021991017,0.00041516824,0.0015484416,0.11589429,0.07301198,0.016452083,0.016709926,0.76948786],"study_design_scores_gemma":[0.00004277332,0.00017659727,0.0008098837,0.000025915437,0.00006422579,0.0002375472,0.00025091157,0.96824205,0.010524835,0.01440401,0.0051739607,0.00004727942],"about_ca_topic_score_codex":0.0063598175,"about_ca_topic_score_gemma":0.010931171,"teacher_disagreement_score":0.0063598175,"about_ca_system_score_codex":0.000602917,"about_ca_system_score_gemma":0.001168925,"threshold_uncertainty_score":0.012645602},"labels":[],"label_agreement":null},{"id":"W4375869140","doi":"10.1109/icassp49357.2023.10096302","title":"SIAST: A Slot Imbalance-Aware Self-Training Scheme for Semi-Supervised Slot Filling","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Natural Science Foundation of China","keywords":"Leverage (statistics); Computer science; Confusion; Training set; Scheme (mathematics); Set (abstract data type); Artificial intelligence; Training (meteorology); Algorithm; Theoretical computer science; Mathematics","score_opus":0.06279102750323409,"score_gpt":0.27805982650759314,"score_spread":0.21526879900435905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4375869140","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01837044,0.00021027321,0.9691133,0.00017549659,0.00013783401,0.00018657178,0.00035642378,0.0102572655,0.0011924206],"genre_scores_gemma":[0.275356,0.00014901397,0.71099925,0.00052035,0.00024906202,0.0008393491,0.0029850353,0.001494658,0.007407293],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99753594,0.00094476994,0.0001435627,0.0007320828,0.00043480826,0.00020893011],"domain_scores_gemma":[0.9947241,0.0023128388,0.00031262983,0.0014764102,0.00089102256,0.0002829797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040814374,0.0013844294,0.0016606916,0.0013310353,0.0012630996,0.0013884783,0.0044586607,0.0022117293,0.004400514],"category_scores_gemma":[0.01112001,0.0008214293,0.0012055282,0.0013091586,0.00128825,0.00406493,0.0035285652,0.0028951643,0.0033564335],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013619951,0.0005169136,0.002931172,0.0002919295,0.00015646199,0.00016153103,0.0012439513,0.072429374,0.027069261,0.0074219075,0.0273393,0.85907626],"study_design_scores_gemma":[0.00007614493,0.00013884854,0.0005114661,0.000018214021,0.000026754553,0.00008992829,0.00011281654,0.9727251,0.012140877,0.009075947,0.00504319,0.00004078586],"about_ca_topic_score_codex":0.002144157,"about_ca_topic_score_gemma":0.005274277,"teacher_disagreement_score":0.0044586607,"about_ca_system_score_codex":0.0008513843,"about_ca_system_score_gemma":0.0015973625,"threshold_uncertainty_score":0.021584988},"labels":[],"label_agreement":null},{"id":"W4376123230","doi":"10.1145/3539618.3591977","title":"SLIM: Sparsified Late Interaction for Multi-Vector Retrieval with Inverted Indexes","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Computer science; Security token; Information retrieval; Inverted index; Vector space model; Artificial intelligence; Search engine indexing","score_opus":0.10816515166389329,"score_gpt":0.3054581386246856,"score_spread":0.19729298696079234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376123230","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064218366,0.00063355325,0.96872485,0.00016467659,0.00015851343,0.00018418768,0.00090607937,0.020200906,0.002605341],"genre_scores_gemma":[0.094203375,0.0005911834,0.88375914,0.00040974724,0.00021109689,0.0005467748,0.0070749954,0.002783365,0.010420329],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986065,0.00022629129,0.00012717042,0.00021475293,0.00065796543,0.00016742307],"domain_scores_gemma":[0.99866354,0.00037677426,0.00010296649,0.0004675401,0.00031108403,0.000078036166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012922562,0.0019608494,0.001869607,0.0016818879,0.00075413106,0.002823519,0.0028856515,0.0013722909,0.017628014],"category_scores_gemma":[0.007296755,0.00066130405,0.0013771428,0.002510779,0.0006851452,0.004917132,0.0040789205,0.0018832496,0.014575324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008366572,0.00021249424,0.0011782885,0.00059769134,0.0001872319,0.00029079744,0.0002683142,0.041467957,0.03259798,0.021134056,0.057637244,0.84359133],"study_design_scores_gemma":[0.00013579281,0.00033663082,0.00063158054,0.000057816942,0.00006491638,0.0004675259,0.00017991461,0.8927528,0.042003483,0.027776536,0.0354753,0.000117658055],"about_ca_topic_score_codex":0.00475217,"about_ca_topic_score_gemma":0.0095911175,"teacher_disagreement_score":0.017628014,"about_ca_system_score_codex":0.0008936066,"about_ca_system_score_gemma":0.0019938394,"threshold_uncertainty_score":0.058971584},"labels":[],"label_agreement":null},{"id":"W4376166979","doi":"10.1145/3539618.3591852","title":"Extracting Complex Named Entities in Legal Documents via Weakly Supervised Object Detection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Dependency (UML); Named-entity recognition; Artificial intelligence; Object (grammar); Information retrieval; Baseline (sea); Information extraction; Object detection; Natural language processing; Data mining; Machine learning; Pattern recognition (psychology); Task (project management)","score_opus":0.03620104001552275,"score_gpt":0.27661519201688234,"score_spread":0.2404141520013596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376166979","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.063219085,0.00081069197,0.91486764,0.00038827397,0.00014647815,0.00020309904,0.0014029131,0.014848015,0.0041138446],"genre_scores_gemma":[0.27840507,0.0005201571,0.7058658,0.0002874685,0.00016790381,0.00021563348,0.0067768916,0.0006600278,0.007100991],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984413,0.00031119192,0.00014539124,0.0006752676,0.00032477343,0.000102058235],"domain_scores_gemma":[0.99439466,0.0022553927,0.0008332564,0.0012089218,0.0011203855,0.00018744345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001478459,0.00092604244,0.0009875958,0.0039287605,0.0008094674,0.002033999,0.0016705722,0.0012880266,0.00227263],"category_scores_gemma":[0.0061628264,0.00039944073,0.0009309119,0.0023783692,0.0008097277,0.0037184146,0.0015936495,0.0013609604,0.004721077],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002425228,0.00033734503,0.010353751,0.0004994983,0.0001263762,0.00078084215,0.0006582345,0.0116558485,0.112041414,0.005794727,0.0153273875,0.84218204],"study_design_scores_gemma":[0.00004018866,0.00021889502,0.01258098,0.000116124844,0.00014993316,0.0013716867,0.0005826861,0.7461513,0.17416374,0.021493059,0.043001954,0.00012938019],"about_ca_topic_score_codex":0.0020945529,"about_ca_topic_score_gemma":0.005492852,"teacher_disagreement_score":0.0039287605,"about_ca_system_score_codex":0.0005200962,"about_ca_system_score_gemma":0.0010587975,"threshold_uncertainty_score":0.007818937},"labels":[],"label_agreement":null},{"id":"W4376619178","doi":"10.1111/cogs.13291","title":"Investigating the Extent to which Distributional Semantic Models Capture a Broad Range of Semantic Relations","year":2023,"lang":"en","type":"article","venue":"Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada; Agencia Estatal de Investigación; Eusko Jaurlaritza; National Science Foundation","keywords":"Semantic similarity; Natural language processing; Computer science; Artificial intelligence; Noun; Pointwise mutual information; Verb; Sentence; Distributional semantics; Similarity (geometry); Noun phrase; Event (particle physics); Mutual information","score_opus":0.051922422714747965,"score_gpt":0.29281857091115643,"score_spread":0.24089614819640848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376619178","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.407551,0.0014802332,0.57076436,0.001799432,0.00013809722,0.0003455972,0.003263628,0.0014887473,0.013169005],"genre_scores_gemma":[0.8861678,0.00046705213,0.10781652,0.00028259284,0.000055455483,0.00029589023,0.00399353,0.0002421397,0.00067902985],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948095,0.0027918282,0.00034758603,0.0013031822,0.00055485585,0.00019297245],"domain_scores_gemma":[0.9776599,0.015455264,0.0015937232,0.0032142664,0.001624065,0.00045282772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012248616,0.0014536594,0.0010907722,0.0039484627,0.0010622238,0.003929115,0.0012971251,0.0012150824,0.0021401248],"category_scores_gemma":[0.040193826,0.00055832136,0.0020388328,0.0033249124,0.0013305779,0.012379008,0.0026494535,0.0019987929,0.0015786091],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015111036,0.0009514907,0.21402915,0.0018184866,0.0023461839,0.00033457438,0.009650167,0.1411542,0.016358571,0.12652586,0.016568907,0.46875134],"study_design_scores_gemma":[0.0000860428,0.00039051176,0.023789829,0.00026879914,0.00029163947,0.0004588365,0.0037842102,0.7325055,0.0030870128,0.2263499,0.008826362,0.00016131118],"about_ca_topic_score_codex":0.005079916,"about_ca_topic_score_gemma":0.008872709,"teacher_disagreement_score":0.012248616,"about_ca_system_score_codex":0.0013983966,"about_ca_system_score_gemma":0.0012683414,"threshold_uncertainty_score":0.06477767},"labels":[],"label_agreement":null},{"id":"W4376632742","doi":"10.48550/arxiv.2305.07393","title":"Prompt Learning to Mitigate Catastrophic Forgetting in Cross-lingual Transfer for Open-domain Dialogue Generation","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Bridging (networking); Computer science; Forgetting; Transfer of learning; Open domain; Natural language processing; Artificial intelligence; Code (set theory); Context (archaeology); Domain (mathematical analysis); Simple (philosophy); Cognitive psychology; Set (abstract data type); Psychology; Programming language; Mathematics","score_opus":0.15415309161359805,"score_gpt":0.24966896185352744,"score_spread":0.09551587023992938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376632742","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10234053,0.0009866204,0.86543816,0.00046289552,0.00037514765,0.00024807005,0.00035407065,0.027445884,0.0023486721],"genre_scores_gemma":[0.79386777,0.0001924443,0.20094067,0.00042798324,0.00010459323,0.00033496498,0.00091178314,0.00097554125,0.002244254],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.997766,0.0011385322,0.000109898465,0.0006390035,0.0002112925,0.000135326],"domain_scores_gemma":[0.9915481,0.005286927,0.00026457294,0.0016422154,0.0008930181,0.0003651648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004158818,0.0014676026,0.00080618163,0.0006188514,0.000582307,0.001087735,0.0018886756,0.0012444444,0.0035184266],"category_scores_gemma":[0.020733193,0.0004356864,0.0005223901,0.0005000548,0.00077665475,0.0034836924,0.003347834,0.0028711217,0.0019801569],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084677554,0.0006866248,0.006677269,0.00075156824,0.0001227661,0.0006001907,0.002248704,0.06145168,0.078550905,0.004954876,0.011278917,0.83182985],"study_design_scores_gemma":[0.00018940594,0.00082064024,0.003445768,0.00007857261,0.000083462575,0.0006123394,0.0007934587,0.8853657,0.07250092,0.023880845,0.012124103,0.00010478161],"about_ca_topic_score_codex":0.0011856277,"about_ca_topic_score_gemma":0.0020597538,"teacher_disagreement_score":0.004158818,"about_ca_system_score_codex":0.0005404959,"about_ca_system_score_gemma":0.0010919985,"threshold_uncertainty_score":0.021994233},"labels":[],"label_agreement":null},{"id":"W4376653681","doi":"10.48550/arxiv.2305.08264","title":"MatSci-NLP: Evaluating Scientific Language Models on Materials Science Language Tasks Using Text-to-Schema Modeling","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Artificial intelligence; Natural language processing; Schema (genetic algorithms); Benchmark (surveying); Question answering; Language model; Natural language understanding; Information retrieval; Natural language","score_opus":0.2573202568645652,"score_gpt":0.28373855231105627,"score_spread":0.026418295446491047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376653681","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44058847,0.012449113,0.23467827,0.0073689315,0.0035591577,0.003536276,0.11435613,0.14141758,0.042046033],"genre_scores_gemma":[0.4903323,0.0016583572,0.19453607,0.0028011429,0.0004307763,0.0034477704,0.29368907,0.0030245387,0.010079939],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959054,0.0016922003,0.00037485964,0.0011154519,0.0006539879,0.00025805668],"domain_scores_gemma":[0.98597383,0.009709422,0.0005474017,0.0018584185,0.0013517627,0.00055913517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066026105,0.0040165186,0.0012852645,0.0028329175,0.0011801112,0.0027322501,0.004526096,0.0039742813,0.011041799],"category_scores_gemma":[0.02449891,0.0008509199,0.0024460857,0.0026778958,0.0015240859,0.005769175,0.0030390846,0.0051863687,0.006661424],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002472468,0.0030621635,0.010837291,0.0048343865,0.0011841565,0.0008235753,0.00062983355,0.40028802,0.0140487095,0.00682723,0.203299,0.35169312],"study_design_scores_gemma":[0.0006065504,0.000953666,0.0040947506,0.0001707103,0.00016407378,0.00024321556,0.00042806222,0.9424859,0.016961688,0.010820048,0.022950452,0.00012083742],"about_ca_topic_score_codex":0.01752574,"about_ca_topic_score_gemma":0.019072026,"teacher_disagreement_score":0.01752574,"about_ca_system_score_codex":0.0028421772,"about_ca_system_score_gemma":0.0040459135,"threshold_uncertainty_score":0.03693849},"labels":[],"label_agreement":null},{"id":"W4376958591","doi":"10.32473/flairs.36.133320","title":"Multi-hop Question Generation without Supporting Fact Information","year":2023,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates; University of Lethbridge","keywords":"Hop (telecommunications); Computer science; Computer network","score_opus":0.22374006884096853,"score_gpt":0.4126515714519567,"score_spread":0.18891150261098816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376958591","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053290226,0.000945919,0.89659816,0.0013841701,0.00027714676,0.0011928516,0.004742992,0.034541346,0.0070271334],"genre_scores_gemma":[0.37099245,0.00036742308,0.60073036,0.0007015272,0.0000998757,0.000553549,0.017585982,0.00086259865,0.008106298],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99841607,0.0005078093,0.00012937945,0.0005387953,0.00032216596,0.0000857363],"domain_scores_gemma":[0.994245,0.0033919788,0.00021798698,0.0012556266,0.00070938916,0.00018002727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026219646,0.0010263734,0.0007611412,0.0012900819,0.00047286254,0.0014121709,0.0025906947,0.0019149318,0.0061421306],"category_scores_gemma":[0.010489047,0.0004162857,0.0014974396,0.0007114537,0.0006633353,0.0041503194,0.0022082366,0.0017138721,0.0026826907],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011901939,0.0009290164,0.007207166,0.0016517635,0.00021517587,0.0015659581,0.0019214234,0.07823758,0.049386796,0.03754582,0.068066105,0.75208306],"study_design_scores_gemma":[0.00020984157,0.000304997,0.0013610027,0.00008534894,0.00012911603,0.0009806197,0.00030797525,0.8649828,0.040512186,0.045398593,0.045672044,0.000055453573],"about_ca_topic_score_codex":0.0026706795,"about_ca_topic_score_gemma":0.004014886,"teacher_disagreement_score":0.0061421306,"about_ca_system_score_codex":0.0008627549,"about_ca_system_score_gemma":0.0011979601,"threshold_uncertainty_score":0.02054751},"labels":[],"label_agreement":null},{"id":"W4377018759","doi":"10.32473/flairs.36.133326","title":"Improving Word Embedding Using Variational Dropout","year":2023,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Carleton University","funders":"","keywords":"Dropout (neural networks); Word (group theory); Computer science; Word embedding; Artificial intelligence; Overfitting; Orthogonality; Natural language processing; Embedding; Curse of dimensionality; Inference; Machine learning; Mathematics; Artificial neural network","score_opus":0.19908569463833295,"score_gpt":0.3984730303623926,"score_spread":0.19938733572405967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377018759","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022836354,0.00040294102,0.974032,0.00027031294,0.00006891049,0.000039966675,0.00012195361,0.0013513643,0.00087619317],"genre_scores_gemma":[0.4892154,0.00095330295,0.49610502,0.0006780776,0.000151392,0.00027357557,0.0022091034,0.00069700653,0.009717136],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918944,0.00027770957,0.000055552508,0.00019818827,0.00019461128,0.000084559106],"domain_scores_gemma":[0.99842864,0.00082538516,0.00012314352,0.00025585789,0.00029952268,0.00006733058],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014866104,0.0012740599,0.0012372748,0.00070537534,0.00040258843,0.00093854696,0.0012441083,0.001312122,0.002306706],"category_scores_gemma":[0.0057675154,0.0005940022,0.0009441224,0.0009599549,0.000928903,0.0037140714,0.0019343136,0.0021781859,0.0011792706],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025782615,0.00027009562,0.0018220497,0.00025843218,0.0001593593,0.00016978358,0.00031227415,0.5099939,0.019868396,0.025385816,0.011205678,0.4302964],"study_design_scores_gemma":[0.0000098378505,0.000024419165,0.00008411009,0.000005178446,0.0000073484275,0.000014685429,0.000011239154,0.99167097,0.0019540382,0.005548051,0.000664423,0.000005670918],"about_ca_topic_score_codex":0.006194534,"about_ca_topic_score_gemma":0.009136427,"teacher_disagreement_score":0.006194534,"about_ca_system_score_codex":0.0008946741,"about_ca_system_score_gemma":0.0015840642,"threshold_uncertainty_score":0.012316942},"labels":[],"label_agreement":null},{"id":"W4377200570","doi":"10.1007/978-3-031-32883-1_55","title":"Preliminary Performance Assessment on Ask4Summary’s Reading Methods for Summary Generation","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Python (programming language); Computer science; Reading (process); Syntax; Mathematics education; Artificial intelligence; World Wide Web; Natural language processing; Programming language; Linguistics; Psychology","score_opus":0.06385710022590928,"score_gpt":0.3447172890229146,"score_spread":0.2808601887970053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377200570","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17097546,0.004038863,0.45072892,0.002963213,0.0023941728,0.0026555557,0.022081513,0.29849043,0.04567185],"genre_scores_gemma":[0.23646039,0.00067384186,0.66370356,0.0008129627,0.0003701969,0.0012116384,0.04937153,0.010279763,0.03711614],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9909253,0.004394506,0.0008528072,0.0015036629,0.001890909,0.00043284017],"domain_scores_gemma":[0.9463535,0.036603317,0.0007416216,0.006974732,0.007968231,0.0013584796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009082771,0.0023324813,0.0017731182,0.0042050164,0.0012587317,0.0037538316,0.0038235933,0.0024116405,0.08518436],"category_scores_gemma":[0.05591201,0.00078044744,0.0017751119,0.0026119843,0.0004935953,0.0048445635,0.003697166,0.0020725657,0.034462538],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002318188,0.00071882876,0.0032022996,0.0010133572,0.00039201565,0.0001772423,0.0008154478,0.0072512673,0.007712344,0.0025388217,0.09796869,0.8758915],"study_design_scores_gemma":[0.002391766,0.0019155694,0.011412258,0.0004892268,0.000890512,0.0007747486,0.0034867732,0.73156375,0.071358,0.009728986,0.16566919,0.00031926687],"about_ca_topic_score_codex":0.0124666365,"about_ca_topic_score_gemma":0.015653852,"teacher_disagreement_score":0.08518436,"about_ca_system_score_codex":0.0014570486,"about_ca_system_score_gemma":0.0022200055,"threshold_uncertainty_score":0.2849701},"labels":[],"label_agreement":null},{"id":"W4377224179","doi":"10.1109/mlbdbi58171.2022.00023","title":"Enhancing BERT-based Passage Retriever with Word Recovery Method","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Labrador Retriever; Computer science; Encoder; Paragraph; Natural language processing; Information retrieval; Word (group theory); Artificial intelligence; Entertainment; Matching (statistics); World Wide Web; Linguistics; Medicine","score_opus":0.017349170740257106,"score_gpt":0.24683582696921694,"score_spread":0.22948665622895983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4377224179","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15481588,0.009945476,0.7551307,0.0012967308,0.0006034299,0.0011688783,0.0076048654,0.059052203,0.010381804],"genre_scores_gemma":[0.47385138,0.0026855452,0.46307957,0.0012345457,0.00064974796,0.00067526347,0.029272337,0.0012472733,0.027304275],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988071,0.00027326192,0.00012547335,0.00029289225,0.00038890462,0.00011245299],"domain_scores_gemma":[0.9984092,0.00050376146,0.00012520442,0.00035155765,0.00053219666,0.00007804531],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014992837,0.0018182434,0.0019485126,0.0036175342,0.00055396085,0.000939866,0.001991429,0.0013498488,0.005869065],"category_scores_gemma":[0.0055032163,0.00035611834,0.0012051866,0.0025027145,0.0005014997,0.003244253,0.0011922432,0.001302075,0.007731751],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008978833,0.0006337045,0.003956459,0.0007343561,0.00024138608,0.00063377613,0.00028926396,0.038488317,0.047941145,0.00272772,0.056153074,0.847303],"study_design_scores_gemma":[0.00022755131,0.0007157997,0.0032637976,0.00004947041,0.00027175664,0.0008930515,0.00020657422,0.92539865,0.03912609,0.0033861857,0.026343577,0.00011753172],"about_ca_topic_score_codex":0.010956704,"about_ca_topic_score_gemma":0.011818432,"teacher_disagreement_score":0.010956704,"about_ca_system_score_codex":0.00073314796,"about_ca_system_score_gemma":0.0013772479,"threshold_uncertainty_score":0.021785855},"labels":[],"label_agreement":null},{"id":"W4378232119","doi":"10.1093/jssam/smad015","title":"Automated Classification for Open-Ended Questions with BERT","year":2023,"lang":"en","type":"article","venue":"Journal of Survey Statistics and Methodology","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Western University","funders":"Social Sciences and Humanities Research Council; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Coding (social sciences); Artificial intelligence; Boosting (machine learning); Natural language processing; Machine learning; Language model; Training set; Statistics","score_opus":0.48587528374401967,"score_gpt":0.4643715025409455,"score_spread":0.021503781203074168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378232119","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55544287,0.0009808594,0.41260457,0.0016896646,0.00035074208,0.00080106047,0.0059223785,0.013896588,0.008311239],"genre_scores_gemma":[0.8951675,0.00009474348,0.09531221,0.00020271224,0.000078450925,0.00037581392,0.0061345287,0.00018565111,0.0024483253],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9930335,0.004431301,0.0003845817,0.00093502435,0.00084858853,0.0003670184],"domain_scores_gemma":[0.9385139,0.04560056,0.00287909,0.006273583,0.0057350774,0.0009978196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009821454,0.0009717103,0.00075563666,0.0024421075,0.0004476257,0.0012810187,0.0015218678,0.001466596,0.0042967983],"category_scores_gemma":[0.05542549,0.00033721264,0.00065036135,0.0015656651,0.0006401632,0.002610463,0.0016496527,0.002226998,0.0036219745],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019803115,0.0010556334,0.09056731,0.0009987942,0.00014893213,0.000266471,0.0022777927,0.06424219,0.01685926,0.009624635,0.043751307,0.76822746],"study_design_scores_gemma":[0.00006830709,0.00030633147,0.029977754,0.00012002933,0.000031288837,0.00016755113,0.0007839292,0.9298141,0.013790926,0.014504878,0.010371138,0.000063738706],"about_ca_topic_score_codex":0.0020516748,"about_ca_topic_score_gemma":0.002885373,"teacher_disagreement_score":0.009821454,"about_ca_system_score_codex":0.0012091104,"about_ca_system_score_gemma":0.0009440864,"threshold_uncertainty_score":0.051941454},"labels":[],"label_agreement":null},{"id":"W4378465150","doi":"10.48550/arxiv.2305.14929","title":"Aligning Language Models to User Opinions","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Ideology; Demographics; Persona; Set (abstract data type); Public opinion; User group; Computer science; Data science; Social psychology; Public relations; Political science; Psychology; World Wide Web; Human–computer interaction; Sociology; Demography; Politics; Law","score_opus":0.17151504849458474,"score_gpt":0.22202192722909095,"score_spread":0.05050687873450621,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378465150","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22416657,0.0008659,0.75701934,0.001964003,0.00020308598,0.00025623626,0.002125686,0.007449665,0.005949482],"genre_scores_gemma":[0.8681766,0.00029790035,0.124909826,0.0003727919,0.00015130179,0.00025887782,0.0027496237,0.000417449,0.00266562],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99659413,0.002040891,0.00017538472,0.0006162078,0.00041010763,0.0001632634],"domain_scores_gemma":[0.9901155,0.0067026,0.0006208331,0.000951213,0.0013875115,0.0002223646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00411496,0.0009962531,0.0007002588,0.0021222269,0.00040326448,0.0023971242,0.0010759055,0.0011649377,0.0026476823],"category_scores_gemma":[0.02258937,0.00053512596,0.0010477279,0.0013669612,0.00041790554,0.0031810512,0.0014453207,0.0016585878,0.002711445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012907496,0.0007155365,0.047679637,0.00058330863,0.000580277,0.00039916634,0.002756657,0.35435796,0.021616261,0.014436136,0.017556379,0.5380279],"study_design_scores_gemma":[0.000014102983,0.0000725721,0.0018064795,0.000024543247,0.00003872819,0.000036982197,0.00022501253,0.9823492,0.002716541,0.0097897155,0.0029012135,0.00002490874],"about_ca_topic_score_codex":0.0053541437,"about_ca_topic_score_gemma":0.006605179,"teacher_disagreement_score":0.0053541437,"about_ca_system_score_codex":0.00093618105,"about_ca_system_score_gemma":0.00080325844,"threshold_uncertainty_score":0.021762252},"labels":[],"label_agreement":null},{"id":"W4378472501","doi":"10.1038/s41598-023-35482-0","title":"Constructing a disease database and using natural language processing to capture and standardize free text clinical information","year":2023,"lang":"en","type":"article","venue":"Scientific Reports","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Public Health Ontario","funders":"Institute of Health Services and Policy Research; Canadian Institutes of Health Research","keywords":"Computer science; Benchmark (surveying); Infectious disease (medical specialty); Disease; Natural language processing; Protocol (science); Data science; Artificial intelligence; Machine learning; Medicine; Pathology; Alternative medicine","score_opus":0.026548197828176384,"score_gpt":0.3204444185690802,"score_spread":0.2938962207409038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378472501","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049033396,0.0008269033,0.9038112,0.0021666326,0.0002830964,0.0021985956,0.02959539,0.008734561,0.0033501408],"genre_scores_gemma":[0.17464045,0.00064291654,0.7597154,0.00051535474,0.00020774343,0.0017367159,0.060499195,0.00025559898,0.0017865675],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963098,0.0010470382,0.0007218186,0.0010982428,0.00070109,0.00012201104],"domain_scores_gemma":[0.9883704,0.007277076,0.001001898,0.0014774663,0.0016263325,0.00024689786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042214952,0.0009336469,0.00092071097,0.008484377,0.00077193807,0.0027836768,0.001794124,0.0012029232,0.0028283065],"category_scores_gemma":[0.016097534,0.0004975005,0.0015626863,0.0043169074,0.00063369353,0.003335052,0.0023447883,0.0017416517,0.0021873855],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047999533,0.0010107493,0.027850538,0.0022066343,0.00026475266,0.001961303,0.0023217953,0.029297745,0.044396587,0.015117177,0.038735647,0.8363572],"study_design_scores_gemma":[0.00026049945,0.0009118688,0.04697814,0.00063375453,0.0005331936,0.004395212,0.0047967476,0.62537473,0.074605815,0.082751386,0.1584039,0.00035477354],"about_ca_topic_score_codex":0.0056069647,"about_ca_topic_score_gemma":0.0051960982,"teacher_disagreement_score":0.008484377,"about_ca_system_score_codex":0.0011074408,"about_ca_system_score_gemma":0.0029997516,"threshold_uncertainty_score":0.022325695},"labels":[],"label_agreement":null},{"id":"W4378513156","doi":"10.48550/arxiv.2305.14975","title":"Just Ask for Calibration: Strategies for Eliciting Calibrated Confidence Scores from Language Models Fine-Tuned with Human Feedback","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Office of Naval Research; Canadian Institute for Advanced Research; National Science Foundation","keywords":"Calibration; Computer science; Ask price; Confidence interval; Trustworthiness; Machine learning; Artificial intelligence; Deferral; Statistics; Mathematics","score_opus":0.1723460406433187,"score_gpt":0.24015724987433215,"score_spread":0.06781120923101344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378513156","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14039902,0.00060371833,0.8306122,0.0012078771,0.00014440874,0.00046644887,0.0010558959,0.021425394,0.0040850854],"genre_scores_gemma":[0.7737534,0.00010961253,0.22181843,0.0005027824,0.00006440459,0.00045034103,0.0011823373,0.00086643646,0.0012523165],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.989767,0.005933,0.0005437377,0.0020299926,0.001377279,0.0003489662],"domain_scores_gemma":[0.901718,0.07649145,0.004771955,0.010240444,0.0055975267,0.0011805983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014626655,0.0024389036,0.0011349423,0.0017819334,0.0005735882,0.002501918,0.0029636994,0.0026993554,0.0037004852],"category_scores_gemma":[0.14933528,0.0008944247,0.00071059354,0.0009782219,0.0012614302,0.0047622165,0.003143066,0.0042159143,0.0018033818],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024907636,0.0011405091,0.038241368,0.0010297463,0.00064160436,0.0005121114,0.004558253,0.19304936,0.04369497,0.013517403,0.01945367,0.68167025],"study_design_scores_gemma":[0.00022876845,0.00039092771,0.0065618497,0.00011490744,0.00008004279,0.00018302017,0.00046311077,0.93015504,0.031401265,0.02665642,0.003587079,0.00017762282],"about_ca_topic_score_codex":0.0021137,"about_ca_topic_score_gemma":0.0031677054,"teacher_disagreement_score":0.014626655,"about_ca_system_score_codex":0.0011521308,"about_ca_system_score_gemma":0.0013886311,"threshold_uncertainty_score":0.07735407},"labels":[],"label_agreement":null},{"id":"W4378770632","doi":"10.48550/arxiv.2305.17446","title":"Fine-tuning Happens in Tiny Subspaces: Exploring Intrinsic Task-specific Subspaces of Pre-trained Language Models","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fundamental Research Funds for the Central Universities; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China; Canadian Institute for Advanced Research","keywords":"Subspace topology; Linear subspace; Parameterized complexity; Computer science; Task (project management); Redundancy (engineering); Outlier; Perspective (graphical); Artificial intelligence; Process (computing); Machine learning; Algorithm; Mathematics","score_opus":0.19224057057225014,"score_gpt":0.21195956701518684,"score_spread":0.019718996442936704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378770632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33738023,0.0008302644,0.65767175,0.0005224346,0.00006171286,0.00005938271,0.00020957006,0.0014600047,0.0018046413],"genre_scores_gemma":[0.9329182,0.0002459832,0.06473338,0.0001838715,0.0000410416,0.000072294795,0.0005445647,0.00028582787,0.0009749257],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99939907,0.00022651753,0.000025614503,0.00020817167,0.000065559936,0.000075152355],"domain_scores_gemma":[0.9975585,0.0014540593,0.00019129415,0.00055698265,0.00013592369,0.00010316819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014282049,0.00096521265,0.00087138027,0.00037248462,0.00038740187,0.0010597288,0.0008011314,0.00080834544,0.0008511714],"category_scores_gemma":[0.010111312,0.0005466137,0.00072346866,0.00040750805,0.0008974843,0.0021879133,0.0013639526,0.0023153576,0.000512017],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029690858,0.00017341068,0.007452612,0.00016795649,0.00018460772,0.00019902726,0.00048523222,0.8238661,0.031390898,0.006597119,0.0020164615,0.12716965],"study_design_scores_gemma":[0.000010035699,0.00005782233,0.0009153713,0.000009941091,0.000014025951,0.000033988978,0.000052539992,0.9867894,0.0036245233,0.007958894,0.0005188195,0.000014631184],"about_ca_topic_score_codex":0.002497102,"about_ca_topic_score_gemma":0.0030338403,"teacher_disagreement_score":0.002497102,"about_ca_system_score_codex":0.0004530397,"about_ca_system_score_gemma":0.00081010745,"threshold_uncertainty_score":0.00755316},"labels":[],"label_agreement":null},{"id":"W4378876515","doi":"10.1007/978-3-031-34107-6_34","title":"Towards Automatic Evaluation of NLG Tasks Using Conversational Large Language Models","year":2023,"lang":"en","type":"book-chapter","venue":"IFIP advances in information and communication technology","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Natural language generation; Computer science; Artificial intelligence; Natural language processing; Natural language","score_opus":0.04236327399656033,"score_gpt":0.3204837379407636,"score_spread":0.2781204639442033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378876515","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06487924,0.003429144,0.83264065,0.0011625735,0.00064779137,0.00073658297,0.0073569315,0.0789366,0.010210509],"genre_scores_gemma":[0.3346021,0.0009393148,0.6185978,0.00047374447,0.00028325463,0.00083943387,0.031118296,0.00478033,0.008365724],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9909733,0.0053620595,0.00049643696,0.0011840449,0.0015804995,0.0004037169],"domain_scores_gemma":[0.9835986,0.010931369,0.00037036248,0.0015792045,0.0030333232,0.00048706488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006694554,0.0026556903,0.0025096878,0.0027308967,0.0013545041,0.0048193987,0.00265944,0.003010979,0.00854979],"category_scores_gemma":[0.020144438,0.0010624491,0.0016030609,0.0017808117,0.00076240575,0.004714434,0.003432679,0.0026384061,0.009161598],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001911859,0.0007291067,0.0036832865,0.0014710365,0.0004112309,0.00045461932,0.0007014299,0.06185853,0.07025128,0.006986267,0.062566526,0.7889749],"study_design_scores_gemma":[0.000098265235,0.00012352722,0.001058321,0.0000546479,0.00008812681,0.00013556555,0.00029210534,0.9588858,0.02361612,0.007328937,0.00827251,0.000046128916],"about_ca_topic_score_codex":0.008820158,"about_ca_topic_score_gemma":0.011014885,"teacher_disagreement_score":0.008820158,"about_ca_system_score_codex":0.0015198975,"about_ca_system_score_gemma":0.0020008038,"threshold_uncertainty_score":0.035404623},"labels":[],"label_agreement":null},{"id":"W4378942418","doi":"10.48550/arxiv.2305.18486","title":"A Systematic Study and Comprehensive Evaluation of ChatGPT on Benchmark Datasets","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Benchmark (surveying); Computer science; Variety (cybernetics); Artificial intelligence; Strengths and weaknesses; Machine learning; Generative grammar; Data science; Machine translation; Benchmarking; Natural language processing; Psychology","score_opus":0.21752327578016173,"score_gpt":0.2541113034357444,"score_spread":0.03658802765558267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378942418","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3204849,0.02496678,0.28445968,0.007567813,0.0036672915,0.0061469865,0.14943242,0.17290448,0.030369617],"genre_scores_gemma":[0.30236393,0.0035875246,0.26628447,0.0027315072,0.0005956226,0.004682301,0.40604886,0.006240913,0.0074649197],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9816234,0.0106850015,0.0012834816,0.0033609422,0.0025220257,0.00052524795],"domain_scores_gemma":[0.9567909,0.026888708,0.0012329681,0.008656266,0.0051568397,0.0012743506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015221173,0.0040021823,0.0020724859,0.005224337,0.002297186,0.0035377902,0.0058980263,0.0030718688,0.005199947],"category_scores_gemma":[0.06210176,0.0009503012,0.0025733672,0.005201782,0.001558605,0.008449485,0.005171049,0.00534093,0.0058986885],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019073475,0.0021395653,0.026173735,0.010242638,0.0017412434,0.0007427966,0.0020925372,0.08697131,0.010260535,0.0082506435,0.37314078,0.47633678],"study_design_scores_gemma":[0.0008762745,0.0024758864,0.025723455,0.0014074567,0.00069354323,0.0017154887,0.0031705538,0.73706555,0.024161441,0.01657446,0.18568374,0.00045207667],"about_ca_topic_score_codex":0.013912005,"about_ca_topic_score_gemma":0.026130568,"teacher_disagreement_score":0.015221173,"about_ca_system_score_codex":0.003480656,"about_ca_system_score_gemma":0.0035214939,"threshold_uncertainty_score":0.08049822},"labels":[],"label_agreement":null},{"id":"W4379053645","doi":"10.3390/cmsf2023006003","title":"Developing Conversational Agent Using Deep Learning Techniques","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Rimouski","funders":"Centre National pour la Recherche Scientifique et Technique","keywords":"Computer science; Converse; Deep learning; Artificial intelligence; Encoder; Natural language; Recurrent neural network; Architecture; Sequence (biology); Artificial neural network; Natural (archaeology); Natural language processing; Human–computer interaction","score_opus":0.08726983656625664,"score_gpt":0.3092228061348674,"score_spread":0.22195296956861077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379053645","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023994701,0.00017705519,0.9681474,0.00034559402,0.0000670538,0.00013917555,0.00009474515,0.0025372244,0.004497104],"genre_scores_gemma":[0.36162567,0.0002431804,0.6289144,0.0003180838,0.000035760815,0.00035956464,0.00034012517,0.00025782804,0.007905297],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99969137,0.00009317835,0.00002705987,0.00007731161,0.000071779345,0.000039394003],"domain_scores_gemma":[0.99952817,0.00019465793,0.00003398972,0.000053557647,0.00014517632,0.000044528653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008584928,0.0005541537,0.0003966474,0.00028548882,0.00050695695,0.0007328661,0.00080565515,0.0008493051,0.0028939478],"category_scores_gemma":[0.0017355997,0.00042517917,0.00069046684,0.00015824854,0.00043112246,0.0015755382,0.0011463854,0.0014129424,0.0009118887],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018746947,0.0004802661,0.0030320063,0.00034111552,0.00020962334,0.00049513066,0.0012756297,0.521758,0.043395456,0.06036817,0.008754224,0.35970294],"study_design_scores_gemma":[0.000006921729,0.000020821337,0.000050114693,0.0000071265217,0.000011780873,0.000021171605,0.00003071175,0.987485,0.004584674,0.0046645193,0.0031111776,0.0000060105262],"about_ca_topic_score_codex":0.0044477843,"about_ca_topic_score_gemma":0.006532348,"teacher_disagreement_score":0.0044477843,"about_ca_system_score_codex":0.0006646291,"about_ca_system_score_gemma":0.0011533552,"threshold_uncertainty_score":0.009681225},"labels":[],"label_agreement":null},{"id":"W4379087247","doi":"10.48550/arxiv.2305.19466","title":"The Impact of Positional Encoding on Length Generalization in Transformers","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Transformer; Generalization; Pairwise comparison; Embedding; Encoding (memory); Computation; Artificial intelligence; Algorithm; Mathematics; Engineering","score_opus":0.11103095269215581,"score_gpt":0.2262783664026312,"score_spread":0.11524741371047538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379087247","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49214906,0.0013106336,0.48615375,0.00092477695,0.00023774915,0.00020268056,0.000903922,0.008291655,0.009825838],"genre_scores_gemma":[0.9157317,0.0005525692,0.07801243,0.00027708107,0.000036678706,0.00011092067,0.001364665,0.00045412153,0.0034598508],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989446,0.00035608522,0.000098307086,0.00036412443,0.0001461254,0.00009074533],"domain_scores_gemma":[0.9930865,0.0045702844,0.00028578352,0.0014367145,0.00045376856,0.00016695839],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00244333,0.001095643,0.0005181464,0.0005093128,0.00033385595,0.0014657445,0.00149854,0.0007621705,0.0039398153],"category_scores_gemma":[0.016626809,0.00045018728,0.00086383376,0.00037149806,0.00076912146,0.00545878,0.0019992571,0.0023136078,0.001658575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001093006,0.0003321982,0.014771113,0.0005619132,0.00015301949,0.00034283154,0.00086321513,0.19073504,0.044443768,0.016292168,0.005670816,0.72474086],"study_design_scores_gemma":[0.00007205411,0.00057696167,0.0028239677,0.00007177737,0.00011795147,0.00037430925,0.00028403962,0.9328294,0.03493689,0.024578558,0.003283241,0.00005086264],"about_ca_topic_score_codex":0.0032696598,"about_ca_topic_score_gemma":0.0051750024,"teacher_disagreement_score":0.0039398153,"about_ca_system_score_codex":0.0008166922,"about_ca_system_score_gemma":0.0010815426,"threshold_uncertainty_score":0.013179958},"labels":[],"label_agreement":null},{"id":"W4379197269","doi":"10.1186/s12859-023-05350-9","title":"An analysis of entity normalization evaluation biases in specialized domains","year":2023,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Normalization (sociology); Computer science; Data science; Task (project management); Natural language processing; Field (mathematics); Information retrieval; Artificial intelligence; Mathematics","score_opus":0.08364742459170624,"score_gpt":0.3394736568482408,"score_spread":0.25582623225653456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379197269","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76299286,0.030737909,0.16640346,0.005215631,0.0011100721,0.0009924578,0.009217371,0.006625447,0.016704803],"genre_scores_gemma":[0.9261284,0.0014586364,0.052194897,0.0010351487,0.0002990049,0.0007129041,0.015124448,0.0013785661,0.0016680096],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9198045,0.051007845,0.007143087,0.009075655,0.011574699,0.0013941573],"domain_scores_gemma":[0.6308101,0.31106392,0.010063841,0.02025882,0.02647473,0.0013285861],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08764417,0.0016230714,0.0015005034,0.004514329,0.0017818096,0.0037995619,0.001868487,0.0018755766,0.0019607113],"category_scores_gemma":[0.24367984,0.00050048495,0.0010913346,0.005996631,0.0018642118,0.0048815426,0.0033800208,0.0019306632,0.001039725],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049888557,0.00087753835,0.23804381,0.0071172044,0.0032761986,0.0009899318,0.005114113,0.035360042,0.025200296,0.014644761,0.080027275,0.5843599],"study_design_scores_gemma":[0.0007856508,0.002082117,0.3624359,0.0036209363,0.0036076661,0.0051154657,0.004959687,0.3185002,0.13819997,0.05870578,0.10140431,0.00058231247],"about_ca_topic_score_codex":0.0027142896,"about_ca_topic_score_gemma":0.0032236117,"teacher_disagreement_score":0.08764417,"about_ca_system_score_codex":0.0021810543,"about_ca_system_score_gemma":0.0015958461,"threshold_uncertainty_score":0.46351218},"labels":[],"label_agreement":null},{"id":"W4379280031","doi":"10.1007/978-3-031-34344-5_21","title":"BERT for Complex Systematic Review Screening to Support the Future of Medical Research","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Workflow; Task (project management); Systematic review; Encoder; Data science; Transformer; Artificial intelligence; Data mining; Machine learning; Information retrieval; Database; MEDLINE","score_opus":0.13582562624998598,"score_gpt":0.3828264763855156,"score_spread":0.2470008501355296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379280031","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014675414,0.1477541,0.19789147,0.5356375,0.032954235,0.0043233065,0.013734608,0.007491906,0.05874536],"genre_scores_gemma":[0.044964608,0.08813015,0.6540045,0.12614799,0.029376445,0.007624656,0.009515191,0.0025328437,0.037703592],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9182888,0.05895473,0.009940826,0.0026868118,0.009528996,0.00059976784],"domain_scores_gemma":[0.4938555,0.40476102,0.02309066,0.024409166,0.046017706,0.007865977],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.11902428,0.001945192,0.0035802224,0.010839089,0.001160161,0.009309481,0.002852625,0.0061426433,0.066525795],"category_scores_gemma":[0.424285,0.0018151115,0.005271351,0.007088386,0.0017034238,0.013867978,0.005214955,0.0063386993,0.01945984],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004473429,0.000049363698,0.00093008525,0.020001726,0.0019263104,0.00019650275,0.0002807888,0.0013064869,0.0004221232,0.08484508,0.58169156,0.30790263],"study_design_scores_gemma":[0.00058830425,0.00025467042,0.0016697395,0.035422616,0.0026900135,0.00048305627,0.00019884612,0.008968947,0.0005264174,0.41053462,0.5384411,0.00022159165],"about_ca_topic_score_codex":0.0025944565,"about_ca_topic_score_gemma":0.0077354545,"teacher_disagreement_score":0.8809757,"about_ca_system_score_codex":0.0037317222,"about_ca_system_score_gemma":0.018961636,"threshold_uncertainty_score":0.6294681},"labels":[],"label_agreement":null},{"id":"W4379374334","doi":"10.21428/594757db.0b1f96f6","title":"ChartSumm: A Comprehensive Benchmark for Automatic Chart Summarization of Long and Short Summaries","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Automatic summarization; Benchmark (surveying); Chart; Metadata; Data mining; Information retrieval; Artificial intelligence; Natural language processing; World Wide Web","score_opus":0.0721323320600067,"score_gpt":0.2976251434578269,"score_spread":0.22549281139782024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379374334","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10266678,0.010520896,0.08354975,0.0024388323,0.001767077,0.0017696227,0.6430817,0.13790922,0.016296038],"genre_scores_gemma":[0.06894537,0.00119717,0.1201052,0.00028941155,0.0002108179,0.0011219622,0.80186445,0.0022573932,0.004008234],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9959346,0.0013485464,0.0006292584,0.00079833146,0.0010416452,0.00024756684],"domain_scores_gemma":[0.9858132,0.006164976,0.0010312198,0.0027005507,0.0035542513,0.0007357879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039914646,0.0023379005,0.0008125087,0.008545877,0.0014496773,0.0024222936,0.0025842427,0.0020342371,0.007759696],"category_scores_gemma":[0.02750558,0.0003538806,0.0014611649,0.0066390303,0.0006545553,0.0040011634,0.0022732671,0.001528387,0.006576571],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012318755,0.00069789594,0.00745069,0.0057039787,0.0002826054,0.0005115132,0.0009742674,0.020602277,0.011601412,0.005809583,0.6649569,0.28017706],"study_design_scores_gemma":[0.00090386724,0.001405602,0.024862774,0.0010632622,0.00028944598,0.00094507786,0.0025488439,0.24671254,0.050214898,0.016361816,0.65433556,0.00035632015],"about_ca_topic_score_codex":0.011446691,"about_ca_topic_score_gemma":0.017190704,"teacher_disagreement_score":0.011446691,"about_ca_system_score_codex":0.0018973869,"about_ca_system_score_gemma":0.0029518409,"threshold_uncertainty_score":0.025958776},"labels":[],"label_agreement":null},{"id":"W4379378423","doi":"10.21428/594757db.8702fa2f","title":"Few shot learning approaches to essay scoring","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Alliance de recherche numérique du Canada","keywords":"Computer science; Shot (pellet); Artificial intelligence; Machine learning; One shot; Training set; Quality (philosophy)","score_opus":0.2935981719903217,"score_gpt":0.2811861866846646,"score_spread":0.01241198530565707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379378423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026195409,0.0008874523,0.9690698,0.00026250907,0.00008796215,0.00018167785,0.00014965991,0.0012273886,0.0019381198],"genre_scores_gemma":[0.6315541,0.0005275596,0.35715201,0.00029880903,0.00027976625,0.0004067752,0.0010240708,0.00023263051,0.008524293],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971915,0.0012671313,0.00013995435,0.0007179449,0.00051754847,0.0001659101],"domain_scores_gemma":[0.99226636,0.0053287777,0.00040258703,0.0006541822,0.0010342501,0.0003138564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033378303,0.0010597213,0.0014941009,0.0019854864,0.00082127354,0.001479494,0.003347977,0.0017820064,0.004133697],"category_scores_gemma":[0.015113328,0.00058849266,0.0007056123,0.001477392,0.0011216133,0.0027817793,0.0017408719,0.0023975612,0.0011672282],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037324755,0.00051172206,0.003471341,0.00042268462,0.00015991743,0.00014577863,0.00043762982,0.45216274,0.004404088,0.02248305,0.0051955706,0.51023227],"study_design_scores_gemma":[0.000017160055,0.000065359,0.00040949916,0.0000107647065,0.000009394447,0.000033521334,0.000025993479,0.9794302,0.0012960975,0.018012779,0.00067503116,0.000014215785],"about_ca_topic_score_codex":0.0045557376,"about_ca_topic_score_gemma":0.005564273,"teacher_disagreement_score":0.0045557376,"about_ca_system_score_codex":0.0014504588,"about_ca_system_score_gemma":0.0011634069,"threshold_uncertainty_score":0.017652333},"labels":[],"label_agreement":null},{"id":"W4379521537","doi":"10.21428/594757db.62395a61","title":"Towards Improving Text Classification Tasks Based on Knowledge Graphs for Limited Labeled Data","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Knowledge graph; Transformer; Artificial intelligence; Domain knowledge; Machine learning; Language model; Graph; Training set; Labeled data; Natural language processing; Process (computing); Data mining; Theoretical computer science","score_opus":0.13882027318002108,"score_gpt":0.3311601625890719,"score_spread":0.1923398894090508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379521537","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09416654,0.0018719225,0.8752679,0.0011453456,0.00015749247,0.00031025297,0.001914544,0.022176705,0.0029892917],"genre_scores_gemma":[0.54683167,0.00081354176,0.43445757,0.000869255,0.00023696365,0.00043577343,0.0111621935,0.00079548155,0.004397529],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983897,0.00054782646,0.000106318555,0.00054828485,0.00027779976,0.00013014501],"domain_scores_gemma":[0.9899306,0.006820031,0.00051408244,0.0015772186,0.00087337854,0.00028466756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002647855,0.0023053368,0.0014434616,0.003953133,0.00081762346,0.0016463803,0.003024414,0.0024140405,0.0020557344],"category_scores_gemma":[0.0111096855,0.00069528644,0.0019073578,0.0030707258,0.00086502807,0.0075691096,0.0020733953,0.003109184,0.002363911],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006994617,0.0009331952,0.0061893803,0.0004340462,0.00024982245,0.00040086362,0.0003782711,0.31883165,0.015185586,0.0065570157,0.023706554,0.6264341],"study_design_scores_gemma":[0.00002449996,0.000037293798,0.0003374133,0.000013835957,0.000027634873,0.000031944553,0.000045628083,0.9884707,0.0027930285,0.0071672583,0.0010404491,0.000010363739],"about_ca_topic_score_codex":0.011382295,"about_ca_topic_score_gemma":0.01558751,"teacher_disagreement_score":0.011382295,"about_ca_system_score_codex":0.0015464717,"about_ca_system_score_gemma":0.0015902361,"threshold_uncertainty_score":0.022632062},"labels":[],"label_agreement":null},{"id":"W4379523158","doi":"10.21428/594757db.63abb0f0","title":"Multihop Factual Claim Verification Using Natural Language Prompts","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Generalization; Artificial intelligence; Task (project management); Natural language processing; Machine learning; Domain (mathematical analysis); Natural language","score_opus":0.046385850745948594,"score_gpt":0.30902863350594095,"score_spread":0.26264278275999237,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379523158","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1144143,0.0006132799,0.85844606,0.0008229806,0.00021707467,0.0003822127,0.0013958417,0.021273473,0.0024347373],"genre_scores_gemma":[0.68297243,0.0002100303,0.31042767,0.00025022705,0.00012478263,0.0002050824,0.0030659675,0.00033051762,0.0024132554],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9957468,0.0013581217,0.00028243306,0.0014412282,0.0009879189,0.00018347344],"domain_scores_gemma":[0.9708736,0.01927085,0.0025085427,0.0041758437,0.0025398536,0.0006313832],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004639635,0.0011537249,0.0011187883,0.0015757941,0.00077006716,0.0016804785,0.0019860864,0.0018302186,0.0043533263],"category_scores_gemma":[0.036763456,0.0004313822,0.0008879822,0.00082193274,0.0009153857,0.005419618,0.0028061566,0.0029898367,0.0018691525],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017251163,0.0009730908,0.012102623,0.0011689776,0.00012260042,0.0020601405,0.0017358266,0.04955304,0.06869192,0.020410145,0.013880714,0.8275759],"study_design_scores_gemma":[0.0001792489,0.00056977686,0.0051156264,0.00012236847,0.00006371769,0.001229251,0.00060299947,0.8556314,0.053021774,0.0675408,0.015816614,0.0001062809],"about_ca_topic_score_codex":0.00086251815,"about_ca_topic_score_gemma":0.0015720826,"teacher_disagreement_score":0.004639635,"about_ca_system_score_codex":0.0007611278,"about_ca_system_score_gemma":0.0018257613,"threshold_uncertainty_score":0.024537086},"labels":[],"label_agreement":null},{"id":"W4380362514","doi":"10.3758/s13421-023-01433-3","title":"A computational account of item-based directed forgetting for nonwords: Incorporating orthographic representations in MINERVA 2","year":2023,"lang":"en","type":"article","venue":"Memory & Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Psychology; Orthographic projection; Cognitive psychology; Motivated forgetting; Forgetting; Natural language processing; Artificial intelligence; Computer science","score_opus":0.03904420266372061,"score_gpt":0.2966687725944468,"score_spread":0.2576245699307262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380362514","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2869764,0.0007055195,0.6986237,0.0016388232,0.0001410517,0.000053809497,0.00055403693,0.0015864179,0.009720227],"genre_scores_gemma":[0.92798924,0.00027982224,0.06785715,0.000106940424,0.00006025602,0.000051808336,0.00036326176,0.0001938357,0.0030975773],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99956626,0.0001584518,0.000028540222,0.0001281086,0.00006269577,0.000055895827],"domain_scores_gemma":[0.99594986,0.0026542123,0.00019939186,0.0007571855,0.00032770864,0.00011171992],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012353131,0.00037283314,0.00079924916,0.0009047031,0.00067251653,0.0022249583,0.0027739834,0.0009591924,0.0048288787],"category_scores_gemma":[0.0072768973,0.0007182614,0.0012156629,0.00095113774,0.00078678597,0.004123661,0.0014134006,0.001337587,0.0008702451],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006837911,0.00031362107,0.01359376,0.0003476942,0.0005141124,0.0007430978,0.0016669461,0.29093763,0.012543449,0.47639635,0.008857753,0.19340193],"study_design_scores_gemma":[0.000024448052,0.00003005946,0.0009849527,0.000013356626,0.00006367801,0.00019111231,0.000047769732,0.7644468,0.0023413606,0.23066643,0.0011623916,0.000027690248],"about_ca_topic_score_codex":0.0026085614,"about_ca_topic_score_gemma":0.0042576524,"teacher_disagreement_score":0.0048288787,"about_ca_system_score_codex":0.0006990566,"about_ca_system_score_gemma":0.0009721595,"threshold_uncertainty_score":0.01615417},"labels":[],"label_agreement":null},{"id":"W4380433238","doi":"10.1145/3588911","title":"A Universal Question-Answering Platform for Knowledge Graphs","year":2023,"lang":"en","type":"article","venue":"Proceedings of the ACM on Management of Data","topic":"Topic Modeling","field":"Computer Science","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"SPARQL; Computer science; Question answering; RDF; Information retrieval; Set (abstract data type); Named graph; Representation (politics); Query language; Artificial intelligence; Semantic Web; Programming language","score_opus":0.09927151819617334,"score_gpt":0.3177039142849232,"score_spread":0.21843239608874981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380433238","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054054754,0.00065944577,0.8065154,0.001115167,0.000179412,0.0007826047,0.00944581,0.16959305,0.0063035935],"genre_scores_gemma":[0.11557583,0.0008121913,0.8234657,0.0015597888,0.00013969121,0.00094891916,0.043438703,0.006509243,0.007549932],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997767,0.0005851346,0.00024322957,0.00078049506,0.00048313726,0.00014097523],"domain_scores_gemma":[0.99570614,0.0017400681,0.00023304255,0.0015358763,0.0005514247,0.00023346407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029688056,0.0012600434,0.0010026973,0.0030214249,0.0011014758,0.002726241,0.0034335274,0.002045022,0.0174754],"category_scores_gemma":[0.013409506,0.000896132,0.002043104,0.0023792973,0.0012530128,0.008511955,0.006681971,0.0028463844,0.008386753],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008156921,0.0005498072,0.0024371282,0.0022748923,0.00024732857,0.0006966301,0.0020136503,0.031247452,0.02512683,0.14139357,0.21072686,0.58247024],"study_design_scores_gemma":[0.00017058048,0.00015227862,0.0012333704,0.00026921832,0.00011008952,0.00050387514,0.00044575735,0.40484232,0.02046156,0.28075486,0.29092717,0.00012882803],"about_ca_topic_score_codex":0.0067533297,"about_ca_topic_score_gemma":0.009189885,"teacher_disagreement_score":0.0174754,"about_ca_system_score_codex":0.0017054051,"about_ca_system_score_gemma":0.0020825577,"threshold_uncertainty_score":0.05846101},"labels":[],"label_agreement":null},{"id":"W4380687254","doi":"10.48550/arxiv.2306.07471","title":"Resources for Brewing BEIR: Reproducible Reference Models and an Official Leaderboard","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Sophistication; Popularity; Artificial intelligence; Benchmark (surveying); CLARITY; Optimal distinctiveness theory; Publication; Information retrieval; Data science; Political science; Geography; Psychology","score_opus":0.32039482435501465,"score_gpt":0.24445730219155654,"score_spread":0.07593752216345812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380687254","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030938484,0.0054658516,0.24958986,0.0051911073,0.0038254533,0.0028083439,0.19531286,0.46147168,0.045396395],"genre_scores_gemma":[0.0636566,0.0012398069,0.25189134,0.0018096187,0.00042630092,0.0042092055,0.625571,0.036202814,0.014993308],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9750993,0.009736231,0.0024206534,0.0036344114,0.0076138517,0.001495552],"domain_scores_gemma":[0.9278285,0.015113419,0.0020065282,0.03696998,0.015465525,0.0026159796],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025357261,0.004585448,0.003058648,0.008010868,0.0025081665,0.0065655,0.0117494725,0.005982025,0.029665278],"category_scores_gemma":[0.10412341,0.002264028,0.0029810725,0.007931211,0.0018371021,0.011064948,0.010359066,0.005758002,0.053623993],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005839962,0.0007444975,0.0014485252,0.0010542859,0.00018956407,0.00016791868,0.00023587176,0.0081649665,0.0019339197,0.007412361,0.8827067,0.095357336],"study_design_scores_gemma":[0.0025574795,0.001714172,0.005927099,0.0014514347,0.00029328722,0.0009606311,0.0009329345,0.2246513,0.030673148,0.047385782,0.68283415,0.00061857665],"about_ca_topic_score_codex":0.018595671,"about_ca_topic_score_gemma":0.028168587,"teacher_disagreement_score":0.029665278,"about_ca_system_score_codex":0.0032233023,"about_ca_system_score_gemma":0.0057121953,"threshold_uncertainty_score":0.1341036},"labels":[],"label_agreement":null},{"id":"W4380992126","doi":"10.1109/access.2023.3286853","title":"Learning to Generate Popular Headlines","year":2023,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence","score_opus":0.08285575657819608,"score_gpt":0.347523334990131,"score_spread":0.2646675784119349,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380992126","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20088089,0.002437976,0.74456054,0.0008290075,0.00044627252,0.00053871825,0.0048498847,0.03703216,0.008424501],"genre_scores_gemma":[0.64671147,0.0009814227,0.31915706,0.0004436359,0.0006325746,0.0005220852,0.015142631,0.0016056169,0.014803571],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99950707,0.0001114704,0.000032491098,0.00019881643,0.00009780886,0.000052281313],"domain_scores_gemma":[0.99710375,0.0013130332,0.00035141082,0.00032846368,0.0007600898,0.00014327923],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007567266,0.0014942773,0.00053489144,0.0018628449,0.0003903018,0.0008101283,0.0008363182,0.00090016914,0.0037928121],"category_scores_gemma":[0.0055771186,0.00039668713,0.0007227672,0.0010659157,0.0002679913,0.0018419377,0.00052730145,0.0007661296,0.0032572218],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007131502,0.00041305408,0.010471089,0.0006941227,0.00018965412,0.00054742076,0.0006086529,0.06956744,0.038563024,0.0043267906,0.0632431,0.81066245],"study_design_scores_gemma":[0.000116144496,0.00041523212,0.0027924918,0.000039953127,0.00011690974,0.00028275905,0.00019104338,0.9487725,0.027435558,0.0047559817,0.015046188,0.00003511603],"about_ca_topic_score_codex":0.0024571198,"about_ca_topic_score_gemma":0.0064846957,"teacher_disagreement_score":0.0037928121,"about_ca_system_score_codex":0.00054659887,"about_ca_system_score_gemma":0.00054510386,"threshold_uncertainty_score":0.012688279},"labels":[],"label_agreement":null},{"id":"W4380995862","doi":"10.3390/app13127126","title":"An Optimized Approach to Translate Technical Patents from English to Japanese Using Machine Translation Models","year":2023,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs","keywords":"Computer science; Machine translation; Artificial intelligence; Natural language processing; Preprocessor; Evaluation of machine translation; Transformer; Translation (biology); Machine learning; Hyperparameter; BLEU; Example-based machine translation; Machine translation software usability; Engineering","score_opus":0.10246256568746862,"score_gpt":0.30376575361934804,"score_spread":0.20130318793187943,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380995862","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025726506,0.00039688256,0.9668561,0.00028299828,0.000068950234,0.00015231536,0.00057773566,0.0029910924,0.002947455],"genre_scores_gemma":[0.34976748,0.0005374878,0.639823,0.00019944731,0.000090070695,0.00038448768,0.0036511438,0.0004613047,0.005085511],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994029,0.00019477839,0.00005839992,0.00017923764,0.000109366476,0.000055302535],"domain_scores_gemma":[0.9994649,0.00021637495,0.000040858813,0.00009784411,0.00015843785,0.000021480035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00090942206,0.0010400277,0.0007154295,0.0009976398,0.0005246749,0.0010660371,0.00076843385,0.00080450095,0.0022683265],"category_scores_gemma":[0.0023833604,0.0004129345,0.0009550292,0.0015319117,0.00030030668,0.0013481509,0.00073525537,0.0009515968,0.0015621218],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023167237,0.00021560202,0.0018668794,0.00038324823,0.00017135967,0.00027391137,0.0002983518,0.2819573,0.021962944,0.012773114,0.01102062,0.6688449],"study_design_scores_gemma":[0.000031851654,0.00008299628,0.0006588984,0.000014668294,0.000066644505,0.00011778895,0.00008668216,0.977239,0.009163931,0.007337215,0.00517869,0.00002169128],"about_ca_topic_score_codex":0.006452919,"about_ca_topic_score_gemma":0.010519324,"teacher_disagreement_score":0.006452919,"about_ca_system_score_codex":0.0006620365,"about_ca_system_score_gemma":0.0024737394,"threshold_uncertainty_score":0.012830734},"labels":[],"label_agreement":null},{"id":"W4381489754","doi":"10.1145/3594778.3594877","title":"EAGER: Explainable Question Answering Using Knowledge Graphs","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; York University; University of Waterloo; Ontario Tech University","funders":"","keywords":"Question answering; Computer science; Knowledge graph; Modular design; Pipeline (software); Graph; Natural language; Artificial intelligence; Natural language processing; Information retrieval; Theoretical computer science; Programming language","score_opus":0.0500733304563139,"score_gpt":0.3037788247930684,"score_spread":0.2537054943367545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381489754","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026673726,0.00026818542,0.93168145,0.00081766857,0.00009788223,0.0002698314,0.0038861765,0.055905916,0.004405526],"genre_scores_gemma":[0.0851561,0.0006623096,0.88621885,0.00079930935,0.00010562678,0.00047493624,0.016107574,0.004183984,0.0062913704],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99772817,0.000946325,0.00014470206,0.00060490886,0.00045710537,0.000118810574],"domain_scores_gemma":[0.99203,0.0060179234,0.00028543102,0.0010571716,0.00045886743,0.000150509],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026373523,0.0016162874,0.0006046869,0.0029969194,0.00077751605,0.0031716651,0.0025427712,0.0023757096,0.031821094],"category_scores_gemma":[0.015538622,0.0008509982,0.0021529992,0.0014476645,0.0011874636,0.0070716757,0.003832391,0.0028745912,0.009000945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003927108,0.00039885455,0.0022751195,0.0022289706,0.0003034845,0.0011341935,0.0027813124,0.038158886,0.017081989,0.23329069,0.14584056,0.5561133],"study_design_scores_gemma":[0.00016546885,0.00011347384,0.0008287175,0.00032917174,0.00011108734,0.00062757294,0.000648839,0.32163605,0.022187473,0.40384486,0.24936898,0.00013833525],"about_ca_topic_score_codex":0.00301159,"about_ca_topic_score_gemma":0.0065488964,"teacher_disagreement_score":0.031821094,"about_ca_system_score_codex":0.0009888249,"about_ca_system_score_gemma":0.0013049186,"threshold_uncertainty_score":0.10645217},"labels":[],"label_agreement":null},{"id":"W4381551981","doi":"10.48550/arxiv.2306.10414","title":"KEST: Kernel Distance Based Efficient Self-Training for Improving Controllable Text Generation","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Text generation; Fluency; Bottleneck; Generator (circuit theory); Natural language generation; Exploit; Kernel (algebra); Language model; Artificial intelligence; Process (computing); Machine learning; Natural language; Mathematics; Power (physics)","score_opus":0.1131849059198157,"score_gpt":0.2004997923629383,"score_spread":0.0873148864431226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381551981","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02753972,0.00031816115,0.9653747,0.00012717897,0.00006615158,0.00006806329,0.000106850894,0.005284325,0.0011149037],"genre_scores_gemma":[0.58554447,0.00026048144,0.40246952,0.00043091705,0.00010057529,0.00033625079,0.0016569125,0.0015643319,0.007636534],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99903035,0.0003590602,0.000064320535,0.00026141424,0.0002020812,0.00008274372],"domain_scores_gemma":[0.9966899,0.0019876305,0.00021568713,0.0005438579,0.000429119,0.00013379657],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001738919,0.001105157,0.00088944123,0.0008340485,0.00047024782,0.0007261283,0.0019143168,0.0012434509,0.002951685],"category_scores_gemma":[0.007484374,0.00049127935,0.0007807267,0.0006625165,0.0009139851,0.002338235,0.0021365352,0.0019065167,0.0018373032],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041012693,0.00036248707,0.0020654413,0.00024517116,0.00009956765,0.00018518778,0.0004174887,0.41438842,0.023485899,0.010762023,0.008553194,0.53902495],"study_design_scores_gemma":[0.000014320679,0.00004602449,0.00011047736,0.000005536383,0.00000526863,0.000028596502,0.000012791647,0.9923701,0.003688579,0.003075072,0.00063582393,0.0000073355095],"about_ca_topic_score_codex":0.0019242613,"about_ca_topic_score_gemma":0.0033368643,"teacher_disagreement_score":0.002951685,"about_ca_system_score_codex":0.00063444284,"about_ca_system_score_gemma":0.000874636,"threshold_uncertainty_score":0.009874344},"labels":[],"label_agreement":null},{"id":"W4381683870","doi":"10.22214/ijraset.2023.53887","title":"Universal Language Model Fine-Tuning for Text Classification","year":2023,"lang":"en","type":"article","venue":"International Journal for Research in Applied Science and Engineering Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":68,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Artificial intelligence; Language model; Transfer of learning; Categorization; Classifier (UML); Natural language processing; Task (project management); Fine-tuning; Deep learning; Process (computing); Natural language; Machine learning; Programming language","score_opus":0.10473531077589725,"score_gpt":0.3965053319277147,"score_spread":0.29177002115181744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381683870","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035121277,0.0016801457,0.93885094,0.0005349472,0.00057276053,0.00021914368,0.00077014026,0.018815953,0.0034347172],"genre_scores_gemma":[0.5322655,0.0006593855,0.44705322,0.0013610636,0.000525725,0.0005945009,0.005453171,0.0016949731,0.01039244],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983871,0.00038428837,0.0001511331,0.0006023858,0.0002717313,0.00020337822],"domain_scores_gemma":[0.997832,0.0008198584,0.00015884716,0.0005330809,0.00052426895,0.00013203335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020394363,0.0015055062,0.0010372511,0.001770159,0.00073086406,0.0015225345,0.001859555,0.0015175525,0.0049951244],"category_scores_gemma":[0.0067619868,0.00035679014,0.001440148,0.0010767193,0.0006756706,0.0037576978,0.0017726725,0.0026039851,0.0043309364],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036815292,0.0005069746,0.0026996157,0.0003386658,0.00017073203,0.0002286404,0.00029721175,0.07249903,0.052037124,0.0062497663,0.03257485,0.83202934],"study_design_scores_gemma":[0.000038568163,0.000096053715,0.0007746771,0.000039298448,0.000041534717,0.00016380353,0.000085733685,0.9462022,0.028413145,0.014293058,0.00980463,0.000047272555],"about_ca_topic_score_codex":0.0040910393,"about_ca_topic_score_gemma":0.0055911257,"teacher_disagreement_score":0.0049951244,"about_ca_system_score_codex":0.0011316906,"about_ca_system_score_gemma":0.0014032532,"threshold_uncertainty_score":0.0167104},"labels":[],"label_agreement":null},{"id":"W4381686872","doi":"10.1162/tacl_a_00564","title":"Questions Are All You Need to Train a Dense Passage Retriever","year":2023,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"DeepMind","keywords":"Computer science; Initialization; Code (set theory); Task (project management); Set (abstract data type); Encoder; Artificial intelligence; Information retrieval; Scheme (mathematics); Training set; Domain (mathematical analysis); Language model; Natural language processing; Machine learning","score_opus":0.030327127669755687,"score_gpt":0.2833717000899555,"score_spread":0.2530445724201998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381686872","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030926432,0.0009200069,0.915548,0.0010389492,0.00036328015,0.00033195474,0.002370875,0.043448154,0.0050523425],"genre_scores_gemma":[0.36474124,0.00048122235,0.60400504,0.0016328715,0.0003314927,0.0005021151,0.01096152,0.003015699,0.014328825],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987608,0.00036985826,0.000076104254,0.0004994872,0.00019573422,0.00009803634],"domain_scores_gemma":[0.9974956,0.0012442976,0.00007598365,0.00071931875,0.0003563901,0.000108357985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022639849,0.0011000491,0.0011488792,0.0007977537,0.00053327536,0.0017151737,0.0020977736,0.0018206605,0.014852851],"category_scores_gemma":[0.0092214085,0.0006025795,0.0013267644,0.00056898437,0.0009781343,0.004097463,0.0024572439,0.0030813138,0.010779975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008348469,0.00033718836,0.0030570242,0.00084203127,0.00024827124,0.00056660466,0.00095581496,0.07300904,0.052695207,0.016959542,0.06127658,0.7892178],"study_design_scores_gemma":[0.000120086894,0.00027919124,0.0011127241,0.00009616559,0.00012717595,0.0005483529,0.00026739988,0.8834192,0.049927082,0.028149815,0.035870858,0.0000818731],"about_ca_topic_score_codex":0.0044152937,"about_ca_topic_score_gemma":0.004821027,"teacher_disagreement_score":0.014852851,"about_ca_system_score_codex":0.00078221987,"about_ca_system_score_gemma":0.0010504614,"threshold_uncertainty_score":0.049687743},"labels":[],"label_agreement":null},{"id":"W4381714234","doi":"10.48550/arxiv.2306.12245","title":"Bidirectional End-to-End Learning of Retriever-Reader Paradigm for Entity Linking","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Tsinghua Shenzhen International Graduate School; Natural Science Foundation of Guangdong Province; National Natural Science Foundation of China","keywords":"End-to-end principle; Labrador Retriever; Computer science; Task (project management); Reading (process); Entity linking; End user; Pipeline (software); Artificial intelligence; Knowledge base; World Wide Web; Engineering; Linguistics; Programming language; Medicine","score_opus":0.13494737724109787,"score_gpt":0.22163611130385824,"score_spread":0.08668873406276037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381714234","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023154337,0.000922445,0.95665747,0.00055546866,0.0001314981,0.00027997495,0.0007424091,0.01319035,0.0043659825],"genre_scores_gemma":[0.3813453,0.0009092268,0.57588816,0.0016192042,0.00022565815,0.00080777163,0.009108557,0.0012944828,0.028801661],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981998,0.0005467986,0.00009514165,0.00075629045,0.00022271155,0.0001791686],"domain_scores_gemma":[0.99602413,0.0020754903,0.00015918753,0.00088237826,0.00071907806,0.00013975825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00348237,0.0026517236,0.0013699079,0.0017065972,0.001113853,0.0018442878,0.0050735385,0.0037160534,0.008623184],"category_scores_gemma":[0.008810909,0.0009887627,0.0015279673,0.0015681671,0.0011473966,0.005926769,0.0036253317,0.004297482,0.008623901],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006045053,0.0007569623,0.004613437,0.0005382959,0.00023589675,0.000556073,0.0007200214,0.12917814,0.015051475,0.009531518,0.029328115,0.8088856],"study_design_scores_gemma":[0.00004649813,0.00021374153,0.0006718236,0.000041173065,0.00007486758,0.0002227347,0.00014174105,0.95536005,0.017746152,0.018097416,0.0073420634,0.000041789153],"about_ca_topic_score_codex":0.0044600177,"about_ca_topic_score_gemma":0.009589251,"teacher_disagreement_score":0.008623184,"about_ca_system_score_codex":0.00092421257,"about_ca_system_score_gemma":0.0014341713,"threshold_uncertainty_score":0.028847396},"labels":[],"label_agreement":null},{"id":"W4381739142","doi":"10.1007/s41870-023-01342-3","title":"A multilingual translator to SQL with database schema pruning to improve self-attention","year":2023,"lang":"en","type":"article","venue":"International Journal of Information Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Fundação de Amparo à Pesquisa do Estado de São Paulo","keywords":"Computer science; Transformer; SQL; Schema (genetic algorithms); Database schema; Artificial intelligence; Natural language processing; Relational database; Information retrieval; Database; Programming language; Database design","score_opus":0.00883787912670021,"score_gpt":0.2724044069227138,"score_spread":0.2635665277960136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381739142","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016666856,0.00036606356,0.7453628,0.00050915504,0.00041158593,0.00021777776,0.0035391138,0.227525,0.0054015964],"genre_scores_gemma":[0.18673107,0.00038397982,0.7222603,0.001265718,0.00025153425,0.00037254952,0.01683947,0.05398795,0.017907405],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978123,0.0004741361,0.00034281233,0.00067200226,0.000536998,0.00016172504],"domain_scores_gemma":[0.992942,0.00232502,0.00023035695,0.0022811093,0.002035604,0.00018603566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019328301,0.0012741751,0.001103786,0.0016021018,0.000679795,0.0028120417,0.0021843486,0.0009599257,0.021974787],"category_scores_gemma":[0.008801565,0.0010714603,0.0012555526,0.0016647213,0.00042449476,0.004389043,0.003959315,0.002166231,0.012542882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000989407,0.000596757,0.004393516,0.001239069,0.00031814078,0.0008563914,0.0032051357,0.00611403,0.06438523,0.019280383,0.20938498,0.68923694],"study_design_scores_gemma":[0.0003477975,0.00028607322,0.0023598208,0.00025792338,0.0004891034,0.0011582536,0.0015361836,0.3987728,0.21320416,0.034353968,0.34700376,0.0002301226],"about_ca_topic_score_codex":0.0026104131,"about_ca_topic_score_gemma":0.0034883951,"teacher_disagreement_score":0.021974787,"about_ca_system_score_codex":0.0005529897,"about_ca_system_score_gemma":0.0015373863,"threshold_uncertainty_score":0.07351291},"labels":[],"label_agreement":null},{"id":"W4382201660","doi":"10.1145/3582768.3582776","title":"Preventing RNN from Using Sequence Length as a Feature","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Actua","funders":"","keywords":"Regularization (linguistics); Computer science; Recurrent neural network; Artificial intelligence; Deep neural networks; Feature (linguistics); Network topology; Sequence (biology); Feature engineering; Artificial neural network; Deep learning; Machine learning; Pattern recognition (psychology)","score_opus":0.06481663986207008,"score_gpt":0.28484853929514586,"score_spread":0.22003189943307577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382201660","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055250745,0.0009321344,0.9309953,0.0024445774,0.0004333001,0.00012330459,0.0002864362,0.0033647947,0.0061695324],"genre_scores_gemma":[0.76964456,0.0007843009,0.21316777,0.0019668436,0.00029497605,0.0003306273,0.0009192329,0.00091668806,0.0119750025],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99818724,0.0006340133,0.00014270675,0.00044925982,0.00046163864,0.00012514532],"domain_scores_gemma":[0.99449533,0.0028645243,0.00046375763,0.0011799322,0.0008637526,0.00013259625],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035433418,0.0015738814,0.0006950639,0.00055091426,0.0007309898,0.0009939986,0.0014670833,0.0016226788,0.0021164867],"category_scores_gemma":[0.021951208,0.0008469042,0.0005751328,0.0006323725,0.00082680664,0.0030941044,0.0019118141,0.0032606714,0.0027592005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007976517,0.00027677914,0.00888103,0.0005091439,0.00020385995,0.00051356974,0.0005611533,0.33885056,0.06660806,0.046939623,0.024223655,0.51163495],"study_design_scores_gemma":[0.00003419298,0.00010700159,0.0013972197,0.000070711656,0.000044619916,0.000270239,0.00007937023,0.93044525,0.025926286,0.03324049,0.008357432,0.000027177304],"about_ca_topic_score_codex":0.005515313,"about_ca_topic_score_gemma":0.008140667,"teacher_disagreement_score":0.005515313,"about_ca_system_score_codex":0.0010486845,"about_ca_system_score_gemma":0.0012414302,"threshold_uncertainty_score":0.018739223},"labels":[],"label_agreement":null},{"id":"W4382201698","doi":"10.1145/3582768.3582774","title":"Extracting Source Information From News Articles","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Credibility; Recall; Variety (cybernetics); Information source (mathematics); Identification (biology); Information retrieval; Attribution; Precision and recall; Software; Data science; Source document; News media; World Wide Web; Artificial intelligence; Political science; Psychology","score_opus":0.02317384518306734,"score_gpt":0.2196091479832862,"score_spread":0.19643530280021887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382201698","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17465529,0.015826939,0.5614722,0.0027774826,0.0015654832,0.0030833923,0.16982096,0.025444359,0.045353852],"genre_scores_gemma":[0.27949053,0.008946003,0.50883335,0.00026074855,0.0021188145,0.0015462984,0.18616052,0.0021166313,0.010527085],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963505,0.0005098227,0.0005247286,0.00067042,0.0017076442,0.00023680803],"domain_scores_gemma":[0.9719559,0.013950394,0.002408771,0.0020983086,0.009099893,0.0004867882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037457978,0.00188198,0.0012665343,0.042111747,0.001519894,0.0047665895,0.0013375609,0.0014967738,0.004610929],"category_scores_gemma":[0.02620399,0.0009840182,0.0013228006,0.023348726,0.00046632945,0.004851433,0.0024283368,0.0017103521,0.005750399],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006710256,0.00031629016,0.04292688,0.0053437557,0.0004646578,0.0024925717,0.0035447166,0.0043680333,0.023116732,0.010122078,0.06854834,0.8380849],"study_design_scores_gemma":[0.0002725796,0.00046337117,0.15742627,0.0033108762,0.00234037,0.0055086734,0.00889838,0.16124459,0.0871701,0.069365196,0.50339985,0.0005997761],"about_ca_topic_score_codex":0.003929414,"about_ca_topic_score_gemma":0.0043813777,"teacher_disagreement_score":0.042111747,"about_ca_system_score_codex":0.0008974248,"about_ca_system_score_gemma":0.0021055308,"threshold_uncertainty_score":0.019809902},"labels":[],"label_agreement":null},{"id":"W4382202705","doi":"10.1609/aaai.v37i11.26486","title":"Adversarial Word Dilution as Text Data Augmentation in Low-Resource Regime","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment","keywords":"Benchmark (surveying); Word (group theory); Computer science; Mixing (physics); Embedding; Resource (disambiguation); Artificial intelligence; Adversarial system; Process (computing); Natural language processing; Word embedding; Class (philosophy); Machine learning; Data mining; Mathematics","score_opus":0.1281727623387476,"score_gpt":0.3286053743921488,"score_spread":0.2004326120534012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382202705","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05119413,0.00081801,0.9414651,0.00084053655,0.00019519683,0.00019161297,0.00033423537,0.0017473426,0.0032139395],"genre_scores_gemma":[0.7924888,0.00046455237,0.1959961,0.0010047633,0.00024657467,0.00060775864,0.0011564159,0.0002991594,0.0077358824],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987601,0.000542925,0.00007218083,0.00029204428,0.00023692749,0.00009591286],"domain_scores_gemma":[0.99481046,0.0036473526,0.00035137616,0.000707453,0.00035597538,0.00012733018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023987794,0.0017123977,0.0012773868,0.00052997,0.00045732528,0.00088493427,0.0017061408,0.0014201291,0.0029523617],"category_scores_gemma":[0.009880307,0.0005109532,0.00085596426,0.0005814154,0.0021022232,0.0027433294,0.0028486252,0.00274975,0.0012582876],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000766985,0.00027744548,0.0024711355,0.00040596706,0.0000991047,0.00040327464,0.0002977972,0.7521851,0.023953997,0.02793911,0.009521735,0.18167837],"study_design_scores_gemma":[0.000024581557,0.00009596217,0.00015646858,0.000020959234,0.000014302269,0.00007289939,0.000017541463,0.9768797,0.006427346,0.014651978,0.0016227227,0.000015501324],"about_ca_topic_score_codex":0.0007427991,"about_ca_topic_score_gemma":0.000949601,"teacher_disagreement_score":0.0029523617,"about_ca_system_score_codex":0.0006811292,"about_ca_system_score_gemma":0.0005812991,"threshold_uncertainty_score":0.012686074},"labels":[],"label_agreement":null},{"id":"W4382239581","doi":"10.1609/aaai.v37i4.25527","title":"Graphs, Constraints, and Search for the Abstraction and Reasoning Corpus","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Theoretical computer science; Artificial intelligence; Graph; Beam search; Machine learning; Search algorithm; Programming language","score_opus":0.11763710036881757,"score_gpt":0.32067466856801335,"score_spread":0.20303756819919577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382239581","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06371859,0.0008383889,0.9167995,0.002259314,0.000071194576,0.00035351358,0.0018349718,0.006127964,0.007996583],"genre_scores_gemma":[0.2552053,0.00041493977,0.73705095,0.00036190965,0.000026217349,0.00026003638,0.0037106688,0.0008476023,0.0021224776],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99670464,0.0016145111,0.00016763636,0.00063396664,0.0007236386,0.00015567982],"domain_scores_gemma":[0.9906887,0.006287751,0.0004790243,0.0018458955,0.00048177224,0.00021693794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032824886,0.00095560396,0.00089241285,0.0018911095,0.0013012926,0.0032104235,0.0023706553,0.0018733452,0.006751785],"category_scores_gemma":[0.018292377,0.0008312973,0.001496776,0.0028143027,0.0033164786,0.009216117,0.0031764235,0.00364196,0.0010776109],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005172952,0.00028006203,0.0027373894,0.0010741886,0.00012316756,0.00039841942,0.0012575053,0.23032235,0.0074664797,0.43961838,0.024728091,0.2914768],"study_design_scores_gemma":[0.000096874675,0.00009016951,0.00057091477,0.000091378075,0.00004320538,0.00014631392,0.00044955005,0.52343374,0.007652839,0.44447187,0.022905113,0.0000480185],"about_ca_topic_score_codex":0.007037593,"about_ca_topic_score_gemma":0.009937043,"teacher_disagreement_score":0.007037593,"about_ca_system_score_codex":0.001992779,"about_ca_system_score_gemma":0.0030976406,"threshold_uncertainty_score":0.022586942},"labels":[],"label_agreement":null},{"id":"W4382318276","doi":"10.1609/aaai.v37i13.26958","title":"Transformer-Based Named Entity Recognition for French Using Adversarial Adaptation to Similar Domain Corpora (Student Abstract)","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Mitacs","keywords":"Named-entity recognition; Computer science; Adversarial system; Transformer; Domain adaptation; Artificial intelligence; Natural language processing; Adaptation (eye); Named entity; Generalization; Engineering; Mathematics","score_opus":0.19473394769154673,"score_gpt":0.33628809076837773,"score_spread":0.141554143076831,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382318276","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.122241296,0.00061944214,0.8657774,0.0007060538,0.0001992966,0.00015225628,0.00079614477,0.0060596764,0.0034483923],"genre_scores_gemma":[0.8105492,0.00034611914,0.17560846,0.00054625905,0.00010908419,0.00014710161,0.0042756377,0.00036543733,0.008052728],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988803,0.00051229575,0.00003972384,0.00035583918,0.000113711605,0.000098077144],"domain_scores_gemma":[0.99818975,0.0010260855,0.00010393099,0.00039365847,0.00023257457,0.00005404162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002505866,0.0009410658,0.0007793705,0.0008407816,0.0004789161,0.0008037991,0.0011463882,0.001056158,0.0015692395],"category_scores_gemma":[0.0036554427,0.0002683472,0.0011090165,0.0008663825,0.0007458573,0.0015118272,0.0012757618,0.0015573307,0.0011040847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048650824,0.00024499654,0.003331398,0.00010552395,0.0002670544,0.0005292177,0.00022501698,0.6133045,0.018841272,0.006393769,0.011985185,0.34428552],"study_design_scores_gemma":[0.0000070066576,0.000033672397,0.0004755265,0.0000044714197,0.000012299003,0.000051163028,0.000016834103,0.9918791,0.0048928126,0.0015736246,0.0010400944,0.000013389305],"about_ca_topic_score_codex":0.009994798,"about_ca_topic_score_gemma":0.0099500455,"teacher_disagreement_score":0.009994798,"about_ca_system_score_codex":0.0008162961,"about_ca_system_score_gemma":0.00047453257,"threshold_uncertainty_score":0.019873261},"labels":[],"label_agreement":null},{"id":"W4382318308","doi":"10.1609/aaai.v37i13.26963","title":"Transformer-Based Multi-Hop Question Generation (Student Abstract)","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Transformer; Language model; Hop (telecommunications); Sentence; Natural language generation; Question answering; Natural language processing; Artificial intelligence; Task (project management); Text generation; Natural language; Telecommunications; Engineering","score_opus":0.16284411487804265,"score_gpt":0.35000643063140036,"score_spread":0.1871623157533577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382318308","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031977568,0.0010749393,0.927321,0.0009965403,0.0002751591,0.0007047419,0.0026156984,0.030333487,0.004700812],"genre_scores_gemma":[0.4612801,0.00034683058,0.52045417,0.00068264094,0.00012152869,0.00041252602,0.008395587,0.00079096784,0.007515592],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998708,0.00046824547,0.000083640516,0.00046365187,0.00019577054,0.00008065968],"domain_scores_gemma":[0.9971174,0.0017589929,0.00010627851,0.0004992652,0.0003930408,0.00012499737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020742537,0.00085485703,0.00071378786,0.0012283395,0.0003883694,0.0010134846,0.0021536096,0.0016843267,0.007247573],"category_scores_gemma":[0.006505536,0.0003814855,0.0016528197,0.00069857616,0.0006512031,0.0024608744,0.001841155,0.001594445,0.0026859078],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009599385,0.00065200153,0.0042382414,0.0008985084,0.00017198829,0.00088285666,0.000829492,0.16077973,0.023512574,0.025671642,0.052830845,0.72857225],"study_design_scores_gemma":[0.00009609888,0.00013771687,0.00044325288,0.000023576898,0.0000451877,0.00029295022,0.00005187138,0.95916057,0.011068818,0.019426936,0.009230082,0.000022924462],"about_ca_topic_score_codex":0.0042088316,"about_ca_topic_score_gemma":0.0058428324,"teacher_disagreement_score":0.007247573,"about_ca_system_score_codex":0.00097020296,"about_ca_system_score_gemma":0.00096833514,"threshold_uncertainty_score":0.0242455},"labels":[],"label_agreement":null},{"id":"W4382318396","doi":"10.1609/aaai.v37i12.26698","title":"Everyone’s Voice Matters: Quantifying Annotation Disagreement Using Demographic Information","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Annotation; Computer science; Demographics; Variance (accounting); Perspective (graphical); Task (project management); Artificial intelligence; Natural language processing; Politeness; Process (computing); Linguistics; Sociology","score_opus":0.16776989053911234,"score_gpt":0.32880023641771317,"score_spread":0.16103034587860082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382318396","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8937303,0.00066015613,0.08831621,0.0011804311,0.00017900899,0.0003497455,0.005396165,0.0010551458,0.0091329245],"genre_scores_gemma":[0.95668584,0.00010801083,0.03134607,0.00026118825,0.00011025045,0.00053736055,0.009488451,0.00017361245,0.0012891833],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9876358,0.0074198907,0.0010129721,0.0018387969,0.0017373089,0.000355175],"domain_scores_gemma":[0.92967683,0.049935315,0.0066021187,0.0046570026,0.008156794,0.00097193645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020297136,0.0009490692,0.000671521,0.0027799124,0.0018555357,0.0018431662,0.0011158984,0.0016502568,0.0011195238],"category_scores_gemma":[0.07717023,0.00036938745,0.0005632956,0.0021728603,0.0009947806,0.0035002755,0.0024815449,0.0014858511,0.0008994301],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022606498,0.0007818303,0.74649644,0.00076637603,0.00042670366,0.0003680426,0.017860522,0.03340018,0.013486309,0.009139551,0.02054522,0.15446821],"study_design_scores_gemma":[0.0002243826,0.00045469817,0.4135848,0.00036947156,0.00028380408,0.00062374375,0.011836442,0.4794059,0.024502853,0.036066487,0.032284614,0.00036284918],"about_ca_topic_score_codex":0.0048962063,"about_ca_topic_score_gemma":0.0061719334,"teacher_disagreement_score":0.020297136,"about_ca_system_score_codex":0.0012737712,"about_ca_system_score_gemma":0.0008947938,"threshold_uncertainty_score":0.10734284},"labels":[],"label_agreement":null},{"id":"W4382318552","doi":"10.1609/aaai.v37i11.26610","title":"Identify Event Causality with Knowledge and Analogy","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Fonds de recherche du Québec – Nature et technologies; Canadian Institute for Advanced Research","keywords":"Analogy; Causality (physics); Computer science; Generalizability theory; Event (particle physics); Identification (biology); Artificial intelligence; Benchmark (surveying); Natural language processing; Sample (material); Machine learning; Data science; Epistemology; Mathematics; Statistics","score_opus":0.10856056689780202,"score_gpt":0.34402306486904427,"score_spread":0.23546249797124225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382318552","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10616073,0.0053819474,0.85961384,0.0025738312,0.00038624022,0.00048136516,0.004882681,0.0063380343,0.014181287],"genre_scores_gemma":[0.76194465,0.0025099774,0.21507555,0.0007499836,0.00034973616,0.00035790293,0.011274753,0.00025321293,0.0074842884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99898213,0.0002067305,0.00008127431,0.00045530268,0.00020081637,0.00007384036],"domain_scores_gemma":[0.9972524,0.0018392381,0.00025757015,0.000310346,0.00026150554,0.00007890427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012029444,0.0015385814,0.00082490884,0.006205455,0.0008135751,0.0015208864,0.0020215656,0.001963779,0.004201444],"category_scores_gemma":[0.0068874853,0.00054272375,0.0018603278,0.0030146854,0.00091151224,0.0065898965,0.0020727601,0.0023540044,0.0013084979],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006176695,0.000637858,0.020552946,0.00097339024,0.00039432477,0.0018382148,0.00079829985,0.1890496,0.004611815,0.053738806,0.030033218,0.6967539],"study_design_scores_gemma":[0.00005019902,0.0000715953,0.0029040086,0.00012174624,0.00012542738,0.00037526028,0.0001964897,0.87736475,0.002203946,0.10150208,0.015041199,0.000043336786],"about_ca_topic_score_codex":0.007667543,"about_ca_topic_score_gemma":0.010101672,"teacher_disagreement_score":0.007667543,"about_ca_system_score_codex":0.0012677298,"about_ca_system_score_gemma":0.0016384231,"threshold_uncertainty_score":0.015245795},"labels":[],"label_agreement":null},{"id":"W4382318960","doi":"10.1609/aaai.v37i12.26706","title":"Deep Learning on a Healthy Data Diet: Finding Important Examples for Fairness","year":2023,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Counterfactual thinking; Spurious relationship; Computer science; Offensive; Machine learning; Pruning; Artificial intelligence; Odds; Equity (law); Tobit model; Focus (optics); Classifier (UML); Random forest; Psychology; Social psychology; Logistic regression; Mathematics; Operations research","score_opus":0.2653499557703358,"score_gpt":0.365566261370321,"score_spread":0.10021630559998518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382318960","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32347217,0.0012717736,0.65927434,0.008306779,0.00027678793,0.00016431815,0.0007763902,0.00083222246,0.005625227],"genre_scores_gemma":[0.83984417,0.00023475541,0.15396209,0.0016016717,0.00019406047,0.00016365344,0.0011806015,0.00015700006,0.0026620368],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99508166,0.0028238615,0.00021811406,0.0009509493,0.00070286513,0.000222573],"domain_scores_gemma":[0.9798055,0.014019847,0.00087877444,0.0035819102,0.0012533402,0.00046064227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012770016,0.0007706379,0.0012956428,0.0009615591,0.001658178,0.002429106,0.0020867744,0.002009709,0.0022270267],"category_scores_gemma":[0.040299326,0.0005531681,0.0006889307,0.00087350165,0.0026917716,0.004543933,0.0033944792,0.003646202,0.0005139435],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030539152,0.00097653124,0.07938587,0.00055121817,0.00042756414,0.0008941192,0.0018683804,0.2529038,0.0078079454,0.10950506,0.027653165,0.51497245],"study_design_scores_gemma":[0.000068705645,0.000106104206,0.002595392,0.00011694795,0.00004825272,0.00011206111,0.00025463276,0.860797,0.0039428123,0.12676622,0.0051683257,0.000023561393],"about_ca_topic_score_codex":0.0018010477,"about_ca_topic_score_gemma":0.00346343,"teacher_disagreement_score":0.012770016,"about_ca_system_score_codex":0.0011377762,"about_ca_system_score_gemma":0.00139901,"threshold_uncertainty_score":0.0675351},"labels":[],"label_agreement":null},{"id":"W4382449327","doi":"10.1162/tacl_a_00556","title":"Aggretriever: A Simple Approach to Aggregate Textual Representations for Robust Dense Passage Retrieval","year":2023,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Language model; Security token; Aggregate (composite); Exploit; Overhead (engineering); Encoder; Code (set theory); Artificial intelligence; Natural language processing; Set (abstract data type); Machine learning; Programming language","score_opus":0.042005300289154314,"score_gpt":0.28978427921608074,"score_spread":0.24777897892692644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382449327","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020907385,0.00068433036,0.93969494,0.00040428122,0.00023892686,0.00034768294,0.0011119251,0.034074266,0.0025362181],"genre_scores_gemma":[0.26034537,0.00045058088,0.72018397,0.0005412574,0.0002004548,0.00052989833,0.0051058675,0.0022494374,0.010393213],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998899,0.00028981408,0.0000851709,0.00037818064,0.00021558935,0.0001323288],"domain_scores_gemma":[0.9983304,0.0005434905,0.00010835925,0.00061240053,0.00030728863,0.00009797362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018460057,0.0018475861,0.0015563731,0.0024087033,0.0006698607,0.0017106321,0.0036889045,0.001475402,0.0106222695],"category_scores_gemma":[0.0060450057,0.00077383756,0.0014893152,0.0019208593,0.0010552312,0.0046816273,0.003638043,0.002208554,0.007996951],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047324933,0.00043074143,0.0013464831,0.00036496986,0.00015487878,0.00036457714,0.00041739063,0.09544717,0.024709344,0.015883382,0.03080574,0.8296021],"study_design_scores_gemma":[0.00006738585,0.0002031749,0.0002706144,0.000029840874,0.00005821669,0.00017830794,0.00011561298,0.9597753,0.012495643,0.016956598,0.009801117,0.000048238897],"about_ca_topic_score_codex":0.0077317148,"about_ca_topic_score_gemma":0.0131170545,"teacher_disagreement_score":0.0106222695,"about_ca_system_score_codex":0.00091564516,"about_ca_system_score_gemma":0.0016891725,"threshold_uncertainty_score":0.035535038},"labels":[],"label_agreement":null},{"id":"W4382456885","doi":"10.1553/giscience2023_01_s140","title":"Ch(e)atGPT? An Anecdotal Approach Addressing the Impact of ChatGPT on Teaching and Learning GIScience","year":2023,"lang":"en","type":"article","venue":"GI_Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Grading (engineering); Coding (social sciences); Computer science; Process (computing); Field (mathematics); Mathematics education; Psychology; Engineering; Sociology; Social science","score_opus":0.06631188049724021,"score_gpt":0.32361811241590765,"score_spread":0.2573062319186674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382456885","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.531223,0.0025076421,0.1628695,0.037246875,0.0040530507,0.0008603013,0.0009816899,0.0046669687,0.2555911],"genre_scores_gemma":[0.91675067,0.00064058794,0.025768917,0.0053570187,0.0005827232,0.0003549834,0.0002258457,0.0006739511,0.049645312],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.99467283,0.004084617,0.0001360005,0.00023989184,0.0006125041,0.00025415365],"domain_scores_gemma":[0.96647537,0.02570126,0.0013418468,0.002634581,0.0023959747,0.0014509195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064914296,0.00039993302,0.00022699422,0.0009731886,0.0019672043,0.0018221153,0.0008512249,0.0023436262,0.009070085],"category_scores_gemma":[0.03023554,0.000218564,0.00022952902,0.000533068,0.0029198937,0.0035271137,0.0036401106,0.00220594,0.001992814],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090867665,0.00042525056,0.015366544,0.00276733,0.000064401924,0.014140344,0.3826391,0.0018393168,0.034648586,0.09060159,0.12546821,0.33113068],"study_design_scores_gemma":[0.000080464,0.0011500205,0.013115673,0.001588691,0.00005218286,0.00999595,0.10213031,0.007593796,0.025177889,0.013290992,0.82564145,0.0001825233],"about_ca_topic_score_codex":0.0006744952,"about_ca_topic_score_gemma":0.0016134706,"teacher_disagreement_score":0.009070085,"about_ca_system_score_codex":0.0009202333,"about_ca_system_score_gemma":0.00052618206,"threshold_uncertainty_score":0.03433031},"labels":[],"label_agreement":null},{"id":"W4382537648","doi":"10.1007/s42979-023-01920-z","title":"Training Integer-Only Deep Recurrent Neural Networks","year":2023,"lang":"en","type":"article","venue":"SN Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Polytechnique Montréal; Huawei Technologies (Canada)","funders":"Huawei Technologies","keywords":"Computer science; Normalization (sociology); Recurrent neural network; Computation; Quantization (signal processing); Edge device; Artificial neural network; Algorithm; Piecewise linear function; Floating point; Artificial intelligence; Mathematics","score_opus":0.06051638147265736,"score_gpt":0.2882258716313546,"score_spread":0.22770949015869724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382537648","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3178246,0.0015708408,0.66572523,0.0009933413,0.00043781483,0.000074024465,0.0007399556,0.0038858466,0.008748293],"genre_scores_gemma":[0.9155627,0.00033704107,0.075238004,0.00019828063,0.00012310845,0.00007423486,0.0012557468,0.0001530395,0.007057829],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998115,0.00003854204,0.000012260023,0.000055148725,0.000031463518,0.00005111111],"domain_scores_gemma":[0.9993062,0.0003960036,0.000047809423,0.00007337937,0.00013247228,0.00004411355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005522094,0.00064109126,0.00059957843,0.00037157597,0.00020724721,0.0006150822,0.0007824602,0.0008809629,0.0030994716],"category_scores_gemma":[0.0026578293,0.00039150994,0.00042730282,0.00038941525,0.00024258367,0.0013945,0.0007519829,0.0011657564,0.0010410704],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041880843,0.0003072193,0.0039036102,0.00018688322,0.00012678879,0.0001827532,0.00011965717,0.5232613,0.017275935,0.014174635,0.009798424,0.43024394],"study_design_scores_gemma":[0.0000052102046,0.000015611835,0.000078357225,0.0000037943705,0.0000062607355,0.0000044622993,0.0000055332694,0.99729604,0.0007537589,0.001633375,0.00019578257,0.0000017854483],"about_ca_topic_score_codex":0.0035712416,"about_ca_topic_score_gemma":0.008071175,"teacher_disagreement_score":0.0035712416,"about_ca_system_score_codex":0.0004493435,"about_ca_system_score_gemma":0.0006465289,"threshold_uncertainty_score":0.010368764},"labels":[],"label_agreement":null},{"id":"W4382537817","doi":"10.1007/s11042-023-15915-8","title":"A Unified Framework for Slot based Response Generation in a Multimodal Dialogue System","year":2023,"lang":"en","type":"article","venue":"Multimedia Tools and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Utterance; Dialog box; Focus (optics); Natural language generation; Natural language; Context (archaeology); Natural language understanding; Task (project management); Exploit; Dialog system; Artificial intelligence; Encoder; Natural language processing; Human–computer interaction; Question answering; World Wide Web","score_opus":0.06605349449938128,"score_gpt":0.29827780494930506,"score_spread":0.23222431044992378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382537817","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010582178,0.000047576064,0.99476624,0.000074525284,0.00002423641,0.000087185406,0.000092064365,0.0028223447,0.0010276416],"genre_scores_gemma":[0.106062196,0.00015864553,0.8884057,0.000102372134,0.00007436567,0.00048086073,0.00048090913,0.00085484434,0.003380082],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976827,0.0008748047,0.00022369895,0.00043815424,0.000555479,0.00022512002],"domain_scores_gemma":[0.99798113,0.00094938016,0.00008311307,0.00035778186,0.00049152615,0.00013705014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038810703,0.0009199949,0.0014322952,0.0014753683,0.0013887025,0.005053931,0.0032280951,0.00196743,0.011249661],"category_scores_gemma":[0.0061288923,0.00096864474,0.0018045555,0.001212853,0.0013751254,0.0035936248,0.0029677842,0.0020360015,0.0048551387],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064858055,0.00031357477,0.0009661859,0.00052943575,0.00019522713,0.00060339447,0.0033049607,0.11424129,0.033473406,0.5262597,0.016810853,0.30265346],"study_design_scores_gemma":[0.00005007644,0.00007842897,0.00015055167,0.00006136208,0.00008566926,0.00014926073,0.00023671552,0.8730515,0.00970808,0.09559967,0.020752957,0.00007579894],"about_ca_topic_score_codex":0.0058493395,"about_ca_topic_score_gemma":0.0058019217,"teacher_disagreement_score":0.011249661,"about_ca_system_score_codex":0.0013391039,"about_ca_system_score_gemma":0.0023747582,"threshold_uncertainty_score":0.037633896},"labels":[],"label_agreement":null},{"id":"W4382542296","doi":"10.1017/s1351324923000311","title":"Korean named entity recognition based on language-specific features","year":2023,"lang":"en","type":"article","venue":"Natural Language Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Morpheme; Natural language processing; Artificial intelligence; Annotation; Named-entity recognition; Scheme (mathematics); Ambiguity; Segmentation; Syllable; Word (group theory); Speech recognition; Task (project management); Linguistics","score_opus":0.010039043517492463,"score_gpt":0.22310503279439328,"score_spread":0.21306598927690082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382542296","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15574741,0.0008731891,0.83165187,0.00027880032,0.00017315232,0.000106502834,0.0010783841,0.005788933,0.004301824],"genre_scores_gemma":[0.7500292,0.0005328913,0.24234214,0.00008902161,0.000046312773,0.00006321108,0.003723466,0.0001840015,0.0029897],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99938977,0.00016133722,0.00005693959,0.00023238032,0.00010510495,0.000054421802],"domain_scores_gemma":[0.99864703,0.0004862421,0.00012377591,0.00032964445,0.00036990427,0.000043355478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010613514,0.00064784486,0.00066586316,0.0012070679,0.00029259242,0.0010093523,0.00093138206,0.00051834347,0.0016174128],"category_scores_gemma":[0.002474282,0.00018835065,0.00088426546,0.0014801233,0.00025356896,0.003891006,0.00066732406,0.0006798471,0.001289908],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005904056,0.00024981756,0.01550482,0.00033448325,0.00036559647,0.00045071755,0.0002859675,0.15452266,0.06398987,0.008916793,0.0068262178,0.7479627],"study_design_scores_gemma":[0.0000071497557,0.000052760563,0.0040274127,0.00001622231,0.00009007432,0.00013783111,0.00007113719,0.964128,0.026018213,0.0017805266,0.0036344968,0.000036160032],"about_ca_topic_score_codex":0.0035173649,"about_ca_topic_score_gemma":0.0050680996,"teacher_disagreement_score":0.0035173649,"about_ca_system_score_codex":0.00032245374,"about_ca_system_score_gemma":0.000385711,"threshold_uncertainty_score":0.0069937706},"labels":[],"label_agreement":null},{"id":"W4382567310","doi":"10.1007/978-3-031-36336-8_83","title":"How Useful Are Educational Questions Generated by Large Language Models?","year":2023,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Canadian Institute for International Peace and Security; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Quality (philosophy); Domain (mathematical analysis); Mathematics education; Natural language processing; Psychology; Epistemology; Mathematics","score_opus":0.05845123054466652,"score_gpt":0.2962198031426312,"score_spread":0.23776857259796466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382567310","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060523573,0.0032366859,0.7784804,0.035144024,0.0007298935,0.00031435527,0.003116467,0.0065010227,0.111953564],"genre_scores_gemma":[0.7176228,0.0027265486,0.24691612,0.0030236782,0.00081049616,0.00043837744,0.0045580952,0.0012915982,0.022612387],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99363214,0.0046260185,0.00015122523,0.00047514215,0.00094385084,0.00017167044],"domain_scores_gemma":[0.9448494,0.048834853,0.00095425575,0.0024371862,0.0023451026,0.00057931186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005180091,0.0008207376,0.0005983028,0.0014595337,0.0005916091,0.005197005,0.0014679861,0.0017600424,0.015987702],"category_scores_gemma":[0.06333567,0.0005357384,0.0007121948,0.0011691032,0.0013076743,0.015633887,0.0015243262,0.0029944314,0.008263087],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054046843,0.0007012066,0.014113665,0.0009716665,0.00014476525,0.0003474106,0.0038468214,0.030731665,0.00553899,0.1974429,0.10481258,0.64080787],"study_design_scores_gemma":[0.00012680751,0.00018744155,0.004084247,0.0004800508,0.00016999757,0.00043798983,0.004372238,0.25477213,0.010989424,0.57361305,0.1506635,0.000103155544],"about_ca_topic_score_codex":0.0018527507,"about_ca_topic_score_gemma":0.001753236,"teacher_disagreement_score":0.015987702,"about_ca_system_score_codex":0.0009655217,"about_ca_system_score_gemma":0.0010402243,"threshold_uncertainty_score":0.05348414},"labels":[],"label_agreement":null},{"id":"W4382567344","doi":"10.1007/978-3-031-36336-8_50","title":"GPTutor: A ChatGPT-Powered Programming Tool for Code Explanation","year":2023,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Personalization; Code (set theory); Source code; Microsoft Visual Studio; Studio; Programming language; Visual programming language; Extension (predicate logic); Multimedia; Software engineering; Program code; Human–computer interaction; World Wide Web; Software","score_opus":0.07773957927845275,"score_gpt":0.31722444014096063,"score_spread":0.23948486086250786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382567344","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00164737,0.00016563217,0.7691074,0.0002799996,0.00024391415,0.00023283376,0.0033463037,0.19079623,0.03418036],"genre_scores_gemma":[0.051794846,0.0008187686,0.5963214,0.0011032056,0.00032971503,0.0018066163,0.0153338555,0.16882351,0.16366811],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99940884,0.00016837195,0.00003902849,0.00010101281,0.00022284206,0.000059943734],"domain_scores_gemma":[0.99660444,0.0024209472,0.00009528919,0.00037970045,0.00034851232,0.00015114795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009130425,0.0016144748,0.00069571665,0.0016771372,0.0005693799,0.0018161384,0.0022334973,0.0013497942,0.1618225],"category_scores_gemma":[0.0060759266,0.00094381307,0.001075028,0.0013332436,0.00050531083,0.0028536723,0.0023795902,0.0022002566,0.04964366],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002745335,0.00012421074,0.00056106993,0.0009926965,0.000045702243,0.0003928249,0.0011625666,0.0046126326,0.009417151,0.043885563,0.49697548,0.44155568],"study_design_scores_gemma":[0.00021386596,0.00010326429,0.00075540546,0.00044034413,0.000067801215,0.0010178663,0.00022752935,0.094183624,0.021587523,0.06348722,0.81781083,0.00010469278],"about_ca_topic_score_codex":0.0012084909,"about_ca_topic_score_gemma":0.0016148692,"teacher_disagreement_score":0.1618225,"about_ca_system_score_codex":0.00065415463,"about_ca_system_score_gemma":0.00090093113,"threshold_uncertainty_score":0.5413502},"labels":[],"label_agreement":null},{"id":"W4382699275","doi":"10.1142/s2717554523500066","title":"Tapping the Potential of Coherence and Syntactic Features in Neural Models for Automatic Essay Scoring","year":2022,"lang":"en","type":"article","venue":"International Journal of Asian Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Coherence (philosophical gambling strategy); Artificial intelligence; Feature engineering; Embedding; Feature (linguistics); Natural language processing; Artificial neural network; Machine learning; Deep learning; Linguistics; Mathematics","score_opus":0.012380619583067244,"score_gpt":0.27862628987198207,"score_spread":0.2662456702889148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382699275","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33644962,0.0008717943,0.65542036,0.00061027164,0.000096842385,0.00007460483,0.000271289,0.0011023676,0.005102756],"genre_scores_gemma":[0.9571218,0.0002049452,0.039761785,0.00006366532,0.000051314444,0.000050429353,0.00027985082,0.00003641155,0.002429872],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995751,0.00019440734,0.000029032582,0.000092216535,0.00007034242,0.00003897458],"domain_scores_gemma":[0.9987024,0.0006472399,0.00014440723,0.0001262124,0.0003279844,0.000051765364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011483665,0.00063033024,0.0003439757,0.00073025905,0.00018720994,0.0007629618,0.00062656077,0.0004820277,0.0013918888],"category_scores_gemma":[0.0045777024,0.00022241022,0.0003411994,0.00078144606,0.0002819752,0.0023560824,0.0007468573,0.0010297254,0.00046249025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039160353,0.00033151187,0.019197116,0.00019466676,0.00027000936,0.00013652438,0.00041310783,0.2572921,0.019220838,0.009041198,0.0028812482,0.6906301],"study_design_scores_gemma":[0.0000052730525,0.000057335663,0.0020449425,0.000008120217,0.00002392496,0.000018374245,0.000035212484,0.9917625,0.0016591867,0.0039834324,0.0003925385,0.000009178693],"about_ca_topic_score_codex":0.0021221654,"about_ca_topic_score_gemma":0.0057731695,"teacher_disagreement_score":0.0021221654,"about_ca_system_score_codex":0.00036948817,"about_ca_system_score_gemma":0.0004876305,"threshold_uncertainty_score":0.0060732365},"labels":[],"label_agreement":null},{"id":"W4383052081","doi":"10.5281/zenodo.8111952","title":"Automated Domain Modeling with Large Language Models: A Comparative Study","year":2023,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Language model; Data modeling; Modeling language; Natural language processing; Programming language; Software engineering; Software; Mathematics","score_opus":0.07540611999253258,"score_gpt":0.30194007886955976,"score_spread":0.22653395887702718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383052081","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24855863,0.0037863574,0.70063496,0.0013610376,0.000105746396,0.00084148016,0.0031020856,0.021930885,0.019678736],"genre_scores_gemma":[0.65510374,0.0017446886,0.33104724,0.00016454812,0.00004039534,0.00031974874,0.00738399,0.002094512,0.0021012384],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9934198,0.0042352066,0.0003863362,0.000517761,0.0013111182,0.0001298095],"domain_scores_gemma":[0.94814855,0.043526173,0.0009390626,0.005330787,0.0017687771,0.00028661845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008567585,0.0007333223,0.0008242375,0.0026565064,0.0008391712,0.0028889643,0.001934299,0.0010567115,0.003797541],"category_scores_gemma":[0.027550515,0.0005777373,0.0013716952,0.0021642065,0.00070055586,0.005116811,0.0021967,0.0015241215,0.0013746833],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019666504,0.001316044,0.014427093,0.0023234754,0.0005742488,0.0005704509,0.003398179,0.21060112,0.0115718255,0.039955232,0.010317781,0.70297784],"study_design_scores_gemma":[0.00020732796,0.00038110398,0.0048152017,0.00022154192,0.00032011475,0.00044557333,0.00117009,0.9190154,0.015606953,0.020290753,0.03744765,0.000078325415],"about_ca_topic_score_codex":0.007625041,"about_ca_topic_score_gemma":0.007706094,"teacher_disagreement_score":0.008567585,"about_ca_system_score_codex":0.0017394198,"about_ca_system_score_gemma":0.0016870599,"threshold_uncertainty_score":0.04531026},"labels":[],"label_agreement":null},{"id":"W4383093199","doi":"10.20944/preprints202307.0192.v1","title":"Examining the Potential of Generative Language Models for Aviation Safety Analysis: Case Study and Insights using the Aviation Safety Reporting System (ASRS)","year":2023,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"General Fusion (Canada)","funders":"","keywords":"Aviation; Computer science; Aviation safety; Similarity (geometry); Crew; Process (computing); Aviation accident; Risk analysis (engineering); Aeronautics; Artificial intelligence; Engineering; Business","score_opus":0.23805610195435514,"score_gpt":0.36995627335307285,"score_spread":0.1319001713987177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383093199","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6478656,0.0011478073,0.33665252,0.0034278438,0.0001355551,0.00050741987,0.00219291,0.002151872,0.0059184013],"genre_scores_gemma":[0.8635597,0.00032403145,0.13195767,0.00024467974,0.000059248086,0.0002030882,0.0023390707,0.00023373881,0.0010786606],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99135756,0.007004127,0.0002106586,0.00068587833,0.00057123316,0.00017038923],"domain_scores_gemma":[0.9344039,0.05927207,0.0015442922,0.0028815498,0.0014958088,0.00040233438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008415361,0.0010639735,0.00036898727,0.0025369835,0.0006475201,0.002341451,0.0011563865,0.0012318387,0.0013439097],"category_scores_gemma":[0.03379125,0.00038862735,0.00116289,0.0016012181,0.0012288562,0.0028934372,0.0016688324,0.0016484406,0.0005421865],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018229722,0.0015445173,0.108115606,0.00248778,0.0006388773,0.00421191,0.039094906,0.4160134,0.022713048,0.05658799,0.0134326415,0.3333364],"study_design_scores_gemma":[0.00008378813,0.00046065878,0.009276548,0.00022867514,0.000133707,0.0008688992,0.0059069875,0.9363088,0.007788851,0.024424816,0.014416009,0.00010223787],"about_ca_topic_score_codex":0.0069736484,"about_ca_topic_score_gemma":0.009554615,"teacher_disagreement_score":0.008415361,"about_ca_system_score_codex":0.0016683803,"about_ca_system_score_gemma":0.001161712,"threshold_uncertainty_score":0.04450518},"labels":[],"label_agreement":null},{"id":"W4383268293","doi":"10.1007/978-3-031-33617-1_3","title":"Methods for Text Generation in NLP","year":2023,"lang":"en","type":"book-chapter","venue":"SpringerBriefs in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland; Toronto Metropolitan University; McGill University; Dalhousie University","funders":"","keywords":"Generative grammar; Computer science; Artificial intelligence; Natural language processing; Text generation; Generative model; Space (punctuation)","score_opus":0.09547071892575935,"score_gpt":0.34734951002743236,"score_spread":0.251878791101673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383268293","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00039250433,0.0010140849,0.99325776,0.00019435247,0.00012488705,0.00004769321,0.0001737661,0.0014357677,0.0033590656],"genre_scores_gemma":[0.026662616,0.0017408171,0.94979304,0.00018427022,0.0003403607,0.00044612857,0.0012580954,0.0014801947,0.018094433],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99751186,0.0011755629,0.0001756467,0.0003747208,0.00067998277,0.000082231134],"domain_scores_gemma":[0.9935962,0.00502363,0.00011557682,0.0007267609,0.00046671723,0.000071148206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002635316,0.0010537017,0.0011407123,0.0019924415,0.00077506294,0.0027606941,0.0023528142,0.0014901664,0.025313634],"category_scores_gemma":[0.011379387,0.0008243372,0.0014571657,0.0024690349,0.0012382611,0.0036329804,0.0025001087,0.001922888,0.0124040665],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008027105,0.00006124716,0.00032012232,0.0008415981,0.00009495377,0.00013413174,0.00037640322,0.026481723,0.0034464935,0.14245035,0.039231036,0.78648174],"study_design_scores_gemma":[0.00005042973,0.000029133595,0.00026378423,0.00016975166,0.000054023287,0.00025998853,0.00013790198,0.44000512,0.0051803184,0.4732658,0.08054229,0.000041470856],"about_ca_topic_score_codex":0.0015520488,"about_ca_topic_score_gemma":0.0020554732,"teacher_disagreement_score":0.025313634,"about_ca_system_score_codex":0.0007580413,"about_ca_system_score_gemma":0.0007747946,"threshold_uncertainty_score":0.084682584},"labels":[],"label_agreement":null},{"id":"W4383605036","doi":"10.48550/arxiv.2307.03042","title":"Parameter-Efficient Fine-Tuning of LLaMA for the Clinical Domain","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"European Commission; Engineering and Physical Sciences Research Council; National Institute for Health and Care Research; Canadian Institute of Steel Construction; UK Research and Innovation; Nvidia; Accenture; Cisco Systems","keywords":"Adapter (computing); Computer science; Downstream (manufacturing); Domain adaptation; Language model; Domain (mathematical analysis); Set (abstract data type); Artificial intelligence; Machine learning; Programming language; Engineering; Mathematics","score_opus":0.25467036970938967,"score_gpt":0.26332279495574423,"score_spread":0.008652425246354567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383605036","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08523366,0.004411294,0.8755906,0.0013363573,0.00038195538,0.0003505234,0.00070692867,0.02703436,0.0049542347],"genre_scores_gemma":[0.7407476,0.001123507,0.24698707,0.001424597,0.00025167057,0.000749789,0.0022461335,0.001587118,0.0048824684],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984603,0.0006567852,0.00011880097,0.0004404719,0.00017934751,0.00014425593],"domain_scores_gemma":[0.9953701,0.0028680216,0.00024687056,0.00080848444,0.0005461034,0.00016042523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035616679,0.0018575336,0.0012135925,0.0011422579,0.000622758,0.0016013728,0.0017651494,0.0017980421,0.0037519685],"category_scores_gemma":[0.018623173,0.00065382454,0.0013178105,0.00072564377,0.000877922,0.0025028407,0.002295317,0.0034636343,0.0041128895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075433933,0.00048885914,0.006851014,0.00038103916,0.00031560907,0.00021597913,0.00029941727,0.21065709,0.034389753,0.0023461804,0.012925245,0.7303754],"study_design_scores_gemma":[0.000107167245,0.0002030888,0.0025020812,0.00007196501,0.00009666264,0.00028895368,0.00014461778,0.9652982,0.017044008,0.007912976,0.006260819,0.000069497575],"about_ca_topic_score_codex":0.0055789175,"about_ca_topic_score_gemma":0.008309737,"teacher_disagreement_score":0.0055789175,"about_ca_system_score_codex":0.0008681249,"about_ca_system_score_gemma":0.0020026423,"threshold_uncertainty_score":0.018836081},"labels":[],"label_agreement":null},{"id":"W4384158719","doi":"10.32620/reks.2023.2.01","title":"Automatic text summarization based on extractive-abstractive method","year":2023,"lang":"en","type":"article","venue":"RADIOELECTRONIC AND COMPUTER SYSTEMS","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Ranking (information retrieval); Natural language processing; Artificial intelligence; Task (project management); WordNet; Meaning (existential); Information retrieval; Face (sociological concept); Noun; Rank (graph theory); Linguistics; Mathematics","score_opus":0.014697214206594023,"score_gpt":0.25804565950512004,"score_spread":0.243348445298526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384158719","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013272822,0.0016299449,0.9618374,0.00035784722,0.00035122182,0.00051694823,0.0018945587,0.018000029,0.0021392668],"genre_scores_gemma":[0.09802239,0.0016996763,0.8762558,0.0002135464,0.0005118215,0.00063760416,0.012698917,0.0010488331,0.008911452],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990006,0.00015616558,0.00014429931,0.00032547297,0.00029341455,0.00007994488],"domain_scores_gemma":[0.9983463,0.0003547399,0.00023501455,0.00020884893,0.0007866347,0.00006831595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007953623,0.0026075232,0.0012721253,0.0034977251,0.00073840993,0.0014747482,0.0011887465,0.0009143854,0.0042595975],"category_scores_gemma":[0.0027566182,0.00048782202,0.0014721266,0.0019730492,0.00041601397,0.002379083,0.00090695143,0.0011380614,0.0042111548],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028841372,0.00012735225,0.0007541282,0.0008741318,0.000117037765,0.0004718181,0.00030348936,0.0121950805,0.09667228,0.0034025097,0.025172405,0.85962135],"study_design_scores_gemma":[0.000173217,0.00064291613,0.0032961236,0.00013234468,0.0004483159,0.0007628606,0.00053466146,0.68007916,0.2286525,0.012282169,0.0728582,0.00013756275],"about_ca_topic_score_codex":0.0020860676,"about_ca_topic_score_gemma":0.002495525,"teacher_disagreement_score":0.0042595975,"about_ca_system_score_codex":0.0005257572,"about_ca_system_score_gemma":0.0010698544,"threshold_uncertainty_score":0.014249742},"labels":[],"label_agreement":null},{"id":"W4384636956","doi":"10.1145/3539618.3591887","title":"MMEAD: MS MARCO Entity Annotations and Disambiguations","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Global Water Futures; Canada First Research Excellence Fund; Natural Sciences and Engineering Research Council of Canada; Nederlandse Organisatie voor Wetenschappelijk Onderzoek","keywords":"Computer science; Python (programming language); Information retrieval; Resource (disambiguation); Named entity; Precision and recall; World Wide Web; Database; Programming language","score_opus":0.04196631591292097,"score_gpt":0.28302388692743785,"score_spread":0.24105757101451689,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384636956","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004096905,0.0008948049,0.035238363,0.00057534676,0.00052174175,0.00030883387,0.7289186,0.20941779,0.02002769],"genre_scores_gemma":[0.009831889,0.0002647475,0.06720373,0.00032444505,0.00011333422,0.0005443889,0.8974802,0.015926626,0.008310714],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99770504,0.00033817196,0.00022728661,0.0008489157,0.000652178,0.0002285364],"domain_scores_gemma":[0.9962191,0.0007200341,0.00029976925,0.0017108168,0.00071152946,0.00033887321],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002340981,0.002550427,0.0013489639,0.005823371,0.0015489934,0.0033219527,0.0036886204,0.0014107152,0.08605714],"category_scores_gemma":[0.012461958,0.0013105746,0.0017461384,0.006493746,0.00065825,0.0064023985,0.0057409494,0.0026624822,0.08501996],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016332051,0.00004597845,0.0009312638,0.00051795895,0.000049289392,0.00010034718,0.00017238638,0.00057626073,0.0015451576,0.0024774233,0.96541893,0.028001662],"study_design_scores_gemma":[0.00014348149,0.00004133579,0.0021361173,0.00014224983,0.00003642024,0.00020353137,0.00016332246,0.006544765,0.0062907864,0.005073155,0.97913945,0.000085377855],"about_ca_topic_score_codex":0.01621891,"about_ca_topic_score_gemma":0.032202993,"teacher_disagreement_score":0.08605714,"about_ca_system_score_codex":0.0012695907,"about_ca_system_score_gemma":0.002475458,"threshold_uncertainty_score":0.28788984},"labels":[],"label_agreement":null},{"id":"W4384640019","doi":"10.1145/3539618.3591801","title":"A Preference Judgment Tool for Authoritative Assessment","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Preference; Computer science; Information retrieval; Misinformation; Transitive relation; Relevance (law); Pairwise comparison; Salience (neuroscience); Artificial intelligence; Statistics; Mathematics","score_opus":0.13050863119496908,"score_gpt":0.3554437093966155,"score_spread":0.22493507820164643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384640019","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017755851,0.00041032556,0.9138636,0.0004996651,0.0003518756,0.0028951457,0.010487133,0.029540049,0.024196332],"genre_scores_gemma":[0.10021251,0.00021290657,0.8753333,0.00030447068,0.00016347122,0.0038814882,0.008232602,0.0027219926,0.008937173],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9712993,0.012914826,0.0036934451,0.0022675477,0.009155446,0.00066948595],"domain_scores_gemma":[0.85767835,0.08677071,0.006061187,0.012449664,0.034646116,0.002393925],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019618643,0.0018013521,0.0013436377,0.009633852,0.0014742322,0.0051513105,0.0018836118,0.001778831,0.031194814],"category_scores_gemma":[0.13985346,0.000746947,0.0014334837,0.0065894127,0.0007089934,0.006442628,0.0040062405,0.0023243101,0.01373517],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012448692,0.0004961098,0.007138134,0.0018768484,0.00019370845,0.00025915468,0.002693553,0.0057606217,0.009209566,0.022099696,0.12774843,0.8212794],"study_design_scores_gemma":[0.0010089729,0.0011514089,0.019684555,0.0013959064,0.00032845626,0.0015669294,0.0031423836,0.29793957,0.038847923,0.14185162,0.49210292,0.0009793324],"about_ca_topic_score_codex":0.001749602,"about_ca_topic_score_gemma":0.0035094502,"teacher_disagreement_score":0.031194814,"about_ca_system_score_codex":0.0011959602,"about_ca_system_score_gemma":0.0028928074,"threshold_uncertainty_score":0.104357004},"labels":[],"label_agreement":null},{"id":"W4384659744","doi":"10.1145/3539618.3592043","title":"Prompt Learning to Mitigate Catastrophic Forgetting in Cross-lingual Transfer for Open-domain Dialogue Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Bridging (networking); Computer science; Transfer of learning; Forgetting; Open domain; Natural language processing; Artificial intelligence; Context (archaeology); Code (set theory); Domain (mathematical analysis); Simple (philosophy); Programming language; Set (abstract data type); Cognitive psychology; Psychology","score_opus":0.05829388737119521,"score_gpt":0.32415341943941745,"score_spread":0.2658595320682222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384659744","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.101540715,0.0010398492,0.8576046,0.00046768098,0.00044090074,0.00029112023,0.00040649125,0.03536079,0.002847849],"genre_scores_gemma":[0.7734044,0.00019817436,0.22073214,0.00047203572,0.000103690116,0.00036671382,0.0010195038,0.0012158025,0.0024874974],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978898,0.0010488712,0.00010758461,0.00061399414,0.00020728896,0.0001324394],"domain_scores_gemma":[0.9923781,0.004676532,0.00024327142,0.0015214416,0.00083573133,0.00034496267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003876932,0.0015347551,0.0007989153,0.00059089117,0.00057300186,0.0011139372,0.0019077192,0.001221218,0.0041614003],"category_scores_gemma":[0.019455686,0.00043033532,0.0005050886,0.0004574685,0.0007743318,0.003455792,0.0033867722,0.0028020274,0.002287316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008400586,0.00068375823,0.005932909,0.0008133478,0.000115886425,0.0006477074,0.0023457224,0.051827393,0.08803331,0.0046685464,0.012357423,0.831734],"study_design_scores_gemma":[0.00022181692,0.0009274256,0.0036015946,0.0000973654,0.00009338659,0.00076716644,0.0009817777,0.86234665,0.09054978,0.02397613,0.016311804,0.00012513249],"about_ca_topic_score_codex":0.0012165838,"about_ca_topic_score_gemma":0.0021179907,"teacher_disagreement_score":0.0041614003,"about_ca_system_score_codex":0.00053107896,"about_ca_system_score_gemma":0.0010958038,"threshold_uncertainty_score":0.020503461},"labels":[],"label_agreement":null},{"id":"W4384695287","doi":"10.22215/etd/2023-15467","title":"Custom Named Construct Recognition in the Business and Management Literature","year":2023,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Named-entity recognition; Computer science; Construct (python library); Information retrieval; Baseline (sea); Information extraction; Data science; Natural language processing; Artificial intelligence; Sample (material); Task (project management); Engineering","score_opus":0.019102564789087405,"score_gpt":0.2534133707422464,"score_spread":0.234310805953159,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384695287","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23186883,0.041290708,0.60855925,0.0072650425,0.001962498,0.0014447524,0.030453598,0.017993178,0.059162106],"genre_scores_gemma":[0.574691,0.011720086,0.3278668,0.0011255632,0.0012873852,0.00079557736,0.06333413,0.0009590576,0.018220406],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970789,0.0008619419,0.00045698485,0.00093809684,0.0005074794,0.0001565285],"domain_scores_gemma":[0.98667604,0.007830124,0.0014995355,0.0016824267,0.0020157942,0.0002961164],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005154197,0.001010222,0.00074993976,0.020899076,0.001320108,0.004109102,0.001345613,0.0013337365,0.00470854],"category_scores_gemma":[0.016375897,0.00042406196,0.0019495441,0.01464152,0.00091140583,0.00783824,0.0023991263,0.0013050858,0.0044024405],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036983396,0.00028125264,0.04173992,0.0037489156,0.00037109453,0.0017464649,0.0029395104,0.014349869,0.014640322,0.027251344,0.032573678,0.85998785],"study_design_scores_gemma":[0.00007434082,0.000373134,0.107927665,0.0023334357,0.0010260108,0.004372639,0.0046047415,0.20375311,0.03450774,0.08584002,0.55486137,0.00032585295],"about_ca_topic_score_codex":0.0055136126,"about_ca_topic_score_gemma":0.00786411,"teacher_disagreement_score":0.020899076,"about_ca_system_score_codex":0.0014194065,"about_ca_system_score_gemma":0.0020459595,"threshold_uncertainty_score":0.027258277},"labels":[],"label_agreement":null},{"id":"W4384698103","doi":"10.22215/etd/2023-15505","title":"Human Emotion and Sentiment in Natural Language Understanding and Generation using Large Language Models with Limited to No Labeled Data","year":2023,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Sentiment analysis; Artificial intelligence; Natural language processing; Natural language generation; Focus (optics); Domain (mathematical analysis); Context (archaeology); Language model; Natural language","score_opus":0.0990675854318619,"score_gpt":0.3320897956854136,"score_spread":0.2330222102535517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384698103","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058146283,0.00041488608,0.92884374,0.0013112336,0.00015058412,0.00024385059,0.0015246534,0.0063038664,0.0030608533],"genre_scores_gemma":[0.50893044,0.0003299632,0.47882995,0.0005564721,0.00011778201,0.00051845686,0.007263398,0.00063027785,0.002823391],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99682546,0.0017409953,0.00017924949,0.00071064476,0.00043455575,0.00010913479],"domain_scores_gemma":[0.9879012,0.008749936,0.0004939538,0.0015929173,0.0010433941,0.00021861658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004992973,0.00090640126,0.0006049622,0.00077786733,0.00064363406,0.0020705764,0.0017290547,0.0012311115,0.0035624958],"category_scores_gemma":[0.02341434,0.00062831945,0.0014188639,0.0005784861,0.00091174716,0.004821901,0.0018737233,0.0026076771,0.0020349885],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010093774,0.00069411215,0.009393442,0.0006063584,0.0003050791,0.00043397467,0.0012675385,0.3985346,0.028012626,0.05329499,0.023876121,0.48257184],"study_design_scores_gemma":[0.000026249372,0.000040376963,0.0004659127,0.000014101676,0.000017922106,0.00003082643,0.000073732146,0.97314024,0.004601543,0.019482872,0.0020910737,0.000015144739],"about_ca_topic_score_codex":0.005626526,"about_ca_topic_score_gemma":0.010366052,"teacher_disagreement_score":0.005626526,"about_ca_system_score_codex":0.0017294869,"about_ca_system_score_gemma":0.0015821755,"threshold_uncertainty_score":0.026405692},"labels":[],"label_agreement":null},{"id":"W4384823501","doi":"10.1145/3539618.3591805","title":"Tevatron: An Efficient and Flexible Toolkit for Neural Retrieval","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Tevatron; Computer science; Ranking (information retrieval); Pipeline (software); Flexibility (engineering); Implementation; Code (set theory); Generalization; Artificial intelligence; Information retrieval; Machine learning; Software engineering; Programming language; Large Hadron Collider; Particle physics","score_opus":0.05170918939528313,"score_gpt":0.29775165103531553,"score_spread":0.2460424616400324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384823501","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071330103,0.0015440227,0.7240615,0.00045392325,0.00028176376,0.00033946286,0.005990424,0.24966727,0.0105285775],"genre_scores_gemma":[0.10760728,0.0017879396,0.81110877,0.0009402149,0.00014276143,0.0012005558,0.02719743,0.025722345,0.024292642],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9991352,0.00016543844,0.00010212752,0.00019242932,0.0003128265,0.0000921126],"domain_scores_gemma":[0.99901164,0.0003976716,0.000059184673,0.00029297685,0.00018554032,0.000052941705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011466586,0.001502817,0.0010627244,0.0014094918,0.00057468785,0.0023288159,0.0038260063,0.0012452162,0.020950906],"category_scores_gemma":[0.0065295575,0.00096525083,0.0015525536,0.0015513233,0.0005437515,0.004792212,0.0037843753,0.002520899,0.018045928],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008153313,0.00028748144,0.0012407727,0.0017501923,0.00039344642,0.00049970415,0.00042926052,0.08246176,0.023391103,0.036842335,0.2892091,0.5626795],"study_design_scores_gemma":[0.00022979855,0.00018046713,0.00051043165,0.00012825344,0.000079463825,0.00039593023,0.00009352694,0.8370114,0.02015532,0.047090825,0.093989536,0.00013505603],"about_ca_topic_score_codex":0.008665132,"about_ca_topic_score_gemma":0.016205084,"teacher_disagreement_score":0.020950906,"about_ca_system_score_codex":0.0011538569,"about_ca_system_score_gemma":0.0015964244,"threshold_uncertainty_score":0.07008779},"labels":[],"label_agreement":null},{"id":"W4384827425","doi":"10.1145/3539618.3591804","title":"One Stop Shop for Question-Answering Dataset Selection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Question answering; Visualization; Point (geometry); Selection (genetic algorithm); Information retrieval; Field (mathematics); Domain (mathematical analysis); World Wide Web; Data science; Artificial intelligence","score_opus":0.062483577979215395,"score_gpt":0.3178259029697293,"score_spread":0.2553423249905139,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384827425","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012847281,0.002082612,0.34983528,0.0034797583,0.0013033443,0.0017923882,0.14992699,0.46529076,0.0134415105],"genre_scores_gemma":[0.07087901,0.00052228116,0.6236321,0.0021567545,0.0004481128,0.003380416,0.25419253,0.036861207,0.007927574],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9906096,0.0030934412,0.0012080518,0.0023576273,0.0021901908,0.0005410259],"domain_scores_gemma":[0.9671414,0.0151167,0.0011415135,0.010406161,0.0045226663,0.0016715985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014093153,0.002333558,0.0021084945,0.007886069,0.0017014003,0.0057837307,0.002656999,0.002452062,0.038229402],"category_scores_gemma":[0.053383533,0.0012280588,0.002202918,0.006341208,0.0007223919,0.007026525,0.008776367,0.0027382283,0.027645135],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008077122,0.00019314028,0.0059148264,0.0014409624,0.00021349973,0.00035423663,0.0009564489,0.0008227255,0.00975718,0.008694173,0.71294063,0.25790438],"study_design_scores_gemma":[0.0006036567,0.00035528291,0.010860191,0.00049232057,0.0001237764,0.0011007236,0.0017321638,0.059593156,0.023551648,0.040463563,0.8608342,0.00028939024],"about_ca_topic_score_codex":0.0029886812,"about_ca_topic_score_gemma":0.006782582,"teacher_disagreement_score":0.038229402,"about_ca_system_score_codex":0.0011343133,"about_ca_system_score_gemma":0.0023567001,"threshold_uncertainty_score":0.12789011},"labels":[],"label_agreement":null},{"id":"W4384891006","doi":"10.1145/3539618.3594243","title":"Uncertainty Quantification for Text Classification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"Engineering and Physical Sciences Research Council","keywords":"Computer science; Uncertainty quantification; Artificial intelligence; Machine learning; Scalability; Robustness (evolution); Calibration; Context (archaeology); Ensemble learning; Deep learning; Mathematics; Statistics; Database","score_opus":0.13152742126255812,"score_gpt":0.32902778271927,"score_spread":0.1975003614567119,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384891006","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012239193,0.0027980746,0.9914954,0.0010483023,0.00026732867,0.000049630788,0.00032747546,0.00043790185,0.0023519683],"genre_scores_gemma":[0.21931528,0.011231534,0.74676514,0.0024951096,0.0057953456,0.001008402,0.003863722,0.0010910811,0.00843431],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9899586,0.003392558,0.00085652963,0.0017828235,0.0036322041,0.00037731594],"domain_scores_gemma":[0.98024625,0.0144000305,0.0012013279,0.0018420403,0.0020017864,0.00030845267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008355423,0.0018408379,0.0016492767,0.004354389,0.0014718651,0.005448939,0.0023950567,0.002454766,0.00766376],"category_scores_gemma":[0.042807236,0.0008977058,0.0022102294,0.003559296,0.0029111914,0.009816564,0.0045319884,0.0055507906,0.002694433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013181678,0.000060699276,0.0013591857,0.00092003815,0.00019315454,0.00019248134,0.0005153041,0.11495404,0.003416761,0.46009398,0.022938268,0.39522427],"study_design_scores_gemma":[0.000007602708,0.000029846618,0.00043597026,0.00021855245,0.000028481963,0.00013761391,0.00007296676,0.33276427,0.0021345478,0.64615023,0.01797089,0.00004905554],"about_ca_topic_score_codex":0.0018868874,"about_ca_topic_score_gemma":0.0011400637,"teacher_disagreement_score":0.008355423,"about_ca_system_score_codex":0.0033683248,"about_ca_system_score_gemma":0.0016730847,"threshold_uncertainty_score":0.04418826},"labels":[],"label_agreement":null},{"id":"W4385144683","doi":"10.1186/s12911-023-02216-1","title":"Quality indices for topic model selection and evaluation: a literature review and case study","year":2023,"lang":"en","type":"review","venue":"BMC Medical Informatics and Decision Making","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Topic model; Artificial intelligence; Automatic summarization; Natural language processing; Model selection; Statistic; Information retrieval; Cluster analysis; Machine learning; Data mining; Statistics; Mathematics","score_opus":0.2952128063180214,"score_gpt":0.5148011429221888,"score_spread":0.21958833660416743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385144683","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006981282,0.92781734,0.057329774,0.0037524637,0.00031958494,0.0006099647,0.00043222975,0.0002964765,0.002460867],"genre_scores_gemma":[0.17264122,0.6455066,0.17235321,0.0021796941,0.0015335941,0.0025127695,0.0021980316,0.00044965243,0.00062528515],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9435896,0.032757387,0.009592783,0.0035423366,0.010012819,0.00050505716],"domain_scores_gemma":[0.57540995,0.38191387,0.011530015,0.004412035,0.025765268,0.00096880546],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0951167,0.0023015332,0.005015519,0.01594582,0.0013699359,0.0068867444,0.004215716,0.0032515808,0.0019974313],"category_scores_gemma":[0.25039926,0.00127624,0.0051069194,0.016133636,0.0021089802,0.00620774,0.002394635,0.0037982825,0.0006890998],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037796752,0.00021516498,0.012905856,0.029545391,0.0014100359,0.00013834175,0.0012247048,0.007931665,0.00036565526,0.009082538,0.012445303,0.9243575],"study_design_scores_gemma":[0.0008240308,0.0035331512,0.06311484,0.24933288,0.01673631,0.005799481,0.008208917,0.28579065,0.009028933,0.08919491,0.26661843,0.0018174596],"about_ca_topic_score_codex":0.006643047,"about_ca_topic_score_gemma":0.005525142,"teacher_disagreement_score":0.9048833,"about_ca_system_score_codex":0.005008863,"about_ca_system_score_gemma":0.005350334,"threshold_uncertainty_score":0.5030312},"labels":[],"label_agreement":null},{"id":"W4385156310","doi":"10.1101/2023.07.17.549421","title":"Identification and Description of Emotions by Current Large Language Models","year":2023,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Identification (biology); Current (fluid); Linguistics; Psychology; Computer science; Natural language processing; Cognitive psychology; Geology; Philosophy","score_opus":0.03460317368317225,"score_gpt":0.24969427387501478,"score_spread":0.21509110019184252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385156310","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.52232885,0.0005290559,0.46563563,0.0011247306,0.000109033055,0.00031027544,0.0016566892,0.0027497034,0.0055560367],"genre_scores_gemma":[0.9266869,0.00014926348,0.07063289,0.00009385889,0.00003563715,0.00015054812,0.001242406,0.00009750199,0.00091092405],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976858,0.0015010302,0.00014603467,0.0003156079,0.00028123963,0.00007028249],"domain_scores_gemma":[0.9903859,0.007231617,0.0007421789,0.00063869444,0.0008534211,0.00014812863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033160045,0.0006455536,0.00032951563,0.0011788808,0.00023445346,0.0020469062,0.0007087046,0.0005274201,0.0017836917],"category_scores_gemma":[0.016637253,0.00017888719,0.000770375,0.000587335,0.00041958233,0.0016580826,0.0008913781,0.0008027553,0.0009083283],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020966663,0.0007439207,0.09625872,0.0009197967,0.0005288803,0.00091648253,0.00955686,0.25068343,0.07479613,0.028114446,0.012026785,0.523358],"study_design_scores_gemma":[0.000033852346,0.00018841156,0.009160426,0.00005548108,0.000075421456,0.00020538156,0.0008157521,0.96624357,0.0075851837,0.0123516135,0.0032376368,0.000047302667],"about_ca_topic_score_codex":0.0017655888,"about_ca_topic_score_gemma":0.0013146814,"teacher_disagreement_score":0.0033160045,"about_ca_system_score_codex":0.0007185336,"about_ca_system_score_gemma":0.0005919712,"threshold_uncertainty_score":0.017536879},"labels":[],"label_agreement":null},{"id":"W4385339878","doi":"10.1007/978-3-031-39841-4_7","title":"Ontology-Driven Parliamentary Analytics: Analysing Political Debates on COVID-19 Impact in Canada","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Ontology; Computer science; Analytics; Context (archaeology); Vocabulary; World Wide Web; Data science; Linguistics","score_opus":0.03953028108057926,"score_gpt":0.2981980090836647,"score_spread":0.2586677280030854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385339878","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90388125,0.0012123622,0.0065916693,0.0036961264,0.000065170076,0.000092202674,0.010873408,0.00026231376,0.07332556],"genre_scores_gemma":[0.98280257,0.0003440506,0.0030785035,0.00010838421,0.00001384377,0.000034158176,0.004726163,0.00013122834,0.008761139],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9985663,0.00026210444,0.000048547372,0.00015409045,0.0005577607,0.00041111367],"domain_scores_gemma":[0.99371547,0.0030508013,0.00042331792,0.00024870227,0.002129822,0.00043191522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019457947,0.00021861547,0.0004025656,0.0037111365,0.00379921,0.0057374695,0.0011668135,0.0007032956,0.004447955],"category_scores_gemma":[0.011884392,0.00023014328,0.00038256063,0.01226623,0.0019223116,0.0023778477,0.0015194432,0.0012008192,0.00054738746],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005843319,0.00021745687,0.37200257,0.0005618019,0.00023599947,0.0006363848,0.07593468,0.034589417,0.002140607,0.24671255,0.088142015,0.17824212],"study_design_scores_gemma":[0.000038269925,0.00003066208,0.45942232,0.00039136526,0.00012900773,0.000084741565,0.1489446,0.10506971,0.0022635704,0.035297997,0.24814886,0.00017889305],"about_ca_topic_score_codex":0.98606616,"about_ca_topic_score_gemma":0.99073166,"teacher_disagreement_score":0.04551197,"about_ca_system_score_codex":0.04551197,"about_ca_system_score_gemma":0.050053507,"threshold_uncertainty_score":0.33021396},"labels":[],"label_agreement":null},{"id":"W4385386988","doi":"10.18280/ria.370320","title":"Evaluation of Short Answers Using Domain Specific Embedding and Siamese Stacked BiLSTM with Contrastive Loss","year":2023,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Domain (mathematical analysis); Embedding; Computer science; Artificial intelligence; Natural language processing; Linguistics; Mathematics; Philosophy","score_opus":0.1025520882371787,"score_gpt":0.33134577830934414,"score_spread":0.22879369007216543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385386988","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7312872,0.0024584236,0.24534117,0.0007667518,0.00051704707,0.00036076366,0.0020187795,0.009339738,0.007910127],"genre_scores_gemma":[0.92396814,0.0003459351,0.06447392,0.0001682025,0.00006270934,0.00014832286,0.004014313,0.00011611833,0.0067023714],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99887246,0.00036525002,0.00009772691,0.00024993668,0.0003249803,0.000089627865],"domain_scores_gemma":[0.99840456,0.00058705173,0.00012878438,0.00013359082,0.00064073165,0.0001053085],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018439849,0.0010965179,0.00065176963,0.00075079105,0.00021870097,0.0009343201,0.0007734207,0.001250967,0.003141096],"category_scores_gemma":[0.0046441336,0.00015062542,0.00047935912,0.00034246317,0.00024886045,0.001364323,0.00079173787,0.0009389807,0.0012590284],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019766546,0.0010468131,0.0128830485,0.0005793154,0.00039068755,0.0004701768,0.00025736034,0.10628705,0.06970417,0.0012734509,0.0127655715,0.7923657],"study_design_scores_gemma":[0.00005406867,0.0013407788,0.0060437014,0.000038275182,0.000061158666,0.00015116604,0.00015303401,0.95270896,0.036393963,0.00096822,0.0020564583,0.000030186187],"about_ca_topic_score_codex":0.0034125827,"about_ca_topic_score_gemma":0.0055683777,"teacher_disagreement_score":0.0034125827,"about_ca_system_score_codex":0.00067865243,"about_ca_system_score_gemma":0.00061786425,"threshold_uncertainty_score":0.01050806},"labels":[],"label_agreement":null},{"id":"W4385400523","doi":"10.18280/ijsse.130310","title":"A Hybrid Text Summarization Approach Using Neural Networks and Metaheuristic Algorithms","year":2023,"lang":"en","type":"article","venue":"International Journal of Safety and Security Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Metaheuristic; Computer science; Artificial neural network; Artificial intelligence; Algorithm; Machine learning","score_opus":0.014499678562923292,"score_gpt":0.23050416823593525,"score_spread":0.21600448967301197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385400523","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018541744,0.00090072857,0.9768673,0.00017095915,0.000063390646,0.00017805849,0.00011803084,0.0013035408,0.0018560704],"genre_scores_gemma":[0.24082425,0.0006943527,0.75178504,0.0002000365,0.00017765252,0.0004494536,0.0007620034,0.00018277409,0.004924384],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99940574,0.00016833466,0.00006542084,0.00016213678,0.00015315971,0.000045113113],"domain_scores_gemma":[0.9993818,0.00025237125,0.00009272946,0.00004511909,0.0002076086,0.00002035124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00088227756,0.0012693662,0.0011714904,0.0025471372,0.00048652894,0.0011399726,0.0011255018,0.0010235836,0.0015806035],"category_scores_gemma":[0.0018091258,0.00038648688,0.0011013236,0.0018227015,0.0002845912,0.0015351914,0.00060124596,0.0006403283,0.00049268245],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021111092,0.00020345075,0.0007165845,0.000456827,0.0003103452,0.00016050886,0.0002323844,0.33614168,0.018608686,0.0059780567,0.0028673634,0.6341131],"study_design_scores_gemma":[0.000024913657,0.00016526414,0.00026415405,0.00001694249,0.0000676907,0.00004050953,0.000058492995,0.9899358,0.004192207,0.0030296159,0.0021892712,0.0000151588465],"about_ca_topic_score_codex":0.0030306792,"about_ca_topic_score_gemma":0.004398008,"teacher_disagreement_score":0.0030306792,"about_ca_system_score_codex":0.00064224616,"about_ca_system_score_gemma":0.00083281077,"threshold_uncertainty_score":0.006026089},"labels":[],"label_agreement":null},{"id":"W4385456320","doi":"10.1145/3611651","title":"Pre-trained Language Models in Biomedical Domain: A Systematic Survey","year":2023,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":166,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Chinese University of Hong Kong","keywords":"Computer science; Biomedical text mining; Terminology; Domain (mathematical analysis); Data science; Artificial intelligence; Taxonomy (biology); Health informatics; Natural language processing; Health care; Text mining","score_opus":0.11915608480073568,"score_gpt":0.3643785811444352,"score_spread":0.24522249634369953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385456320","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006596919,0.8005201,0.17084983,0.008034397,0.0013945069,0.0002736709,0.0018996435,0.0021224974,0.008308328],"genre_scores_gemma":[0.072297834,0.8029664,0.097561345,0.005436184,0.0026112825,0.00060734,0.009911335,0.0006840419,0.007924158],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.99786794,0.000875172,0.00023391075,0.00046064015,0.00047633308,0.00008608099],"domain_scores_gemma":[0.98688656,0.010820453,0.00025915488,0.00049197103,0.0014140842,0.00012768505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004771274,0.0018833015,0.0016007478,0.0029430324,0.00032060684,0.0017474988,0.0025213857,0.0015213774,0.0039810915],"category_scores_gemma":[0.019145118,0.00066929497,0.001698137,0.002537996,0.0007269297,0.0049694073,0.0012923743,0.00242612,0.0043937336],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000994507,0.00013861113,0.0021423995,0.007444938,0.00025488285,0.000107845895,0.00020733297,0.009304167,0.0012183496,0.005934558,0.03132462,0.94182295],"study_design_scores_gemma":[0.000118490236,0.0011687475,0.010283204,0.019454297,0.0017501846,0.0024006702,0.0012666868,0.17567244,0.011360992,0.052656993,0.72349143,0.00037588913],"about_ca_topic_score_codex":0.0045792386,"about_ca_topic_score_gemma":0.0055951066,"teacher_disagreement_score":0.004771274,"about_ca_system_score_codex":0.0010483594,"about_ca_system_score_gemma":0.0034196752,"threshold_uncertainty_score":0.02523321},"labels":[],"label_agreement":null},{"id":"W4385489610","doi":"10.1109/cai54212.2023.00097","title":"Can AI have a personality?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Agreeableness; Conscientiousness; Personality; Neuroticism; Hierarchical structure of the Big Five; Computer science; Big Five personality traits; Psychology; Extraversion and introversion; Artificial intelligence; Social psychology","score_opus":0.04690695920680991,"score_gpt":0.284228868406088,"score_spread":0.23732190919927812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385489610","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9338847,0.0010370332,0.01299382,0.010759124,0.00029503764,0.00007144325,0.00029483403,0.00031582706,0.040348344],"genre_scores_gemma":[0.99562824,0.00021235507,0.0019743945,0.0005029086,0.00005190939,0.000017337758,0.000073931515,0.000020844096,0.001518158],"study_design_codex":"observational","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99809223,0.0010825888,0.000078508376,0.00023022294,0.00032570714,0.00019073252],"domain_scores_gemma":[0.97789997,0.012076032,0.0024246385,0.003001123,0.0019801753,0.0026180402],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048612524,0.00041028045,0.00033082446,0.0008087904,0.0011685748,0.0034176966,0.00036161297,0.0006626586,0.0033287925],"category_scores_gemma":[0.026033353,0.00023750457,0.00037991026,0.00082654937,0.001619282,0.0031936686,0.0010748866,0.001534519,0.0010466166],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005539138,0.00038594016,0.7638757,0.0002702526,0.00038933975,0.0010566086,0.019217672,0.004069078,0.003364988,0.029526183,0.008658622,0.1686316],"study_design_scores_gemma":[0.00009423194,0.0010355331,0.78957295,0.00022450853,0.00029100536,0.0027350448,0.022403553,0.037386872,0.0018667509,0.11401419,0.030133108,0.00024215065],"about_ca_topic_score_codex":0.0027585353,"about_ca_topic_score_gemma":0.0027901982,"teacher_disagreement_score":0.0048612524,"about_ca_system_score_codex":0.0006399473,"about_ca_system_score_gemma":0.0005371952,"threshold_uncertainty_score":0.025709033},"labels":[],"label_agreement":null},{"id":"W4385546024","doi":"10.1016/j.patter.2023.100802","title":"Evaluating progress in automatic chest X-ray radiology report generation","year":2023,"lang":"en","type":"article","venue":"Patterns","topic":"Topic Modeling","field":"Computer Science","cited_by":151,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"National Institute of Biomedical Imaging and Bioengineering; National Institutes of Health; International Business Machines Corporation; National Heart, Lung, and Blood Institute; Gordon and Betty Moore Foundation","keywords":"Workload; Computer science; Metric (unit); Correctness; Selection (genetic algorithm); Artificial intelligence; Medical physics; Machine learning; Radiology; Data mining; Medicine; Algorithm; Engineering","score_opus":0.12598750857622068,"score_gpt":0.37312204709445873,"score_spread":0.24713453851823805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385546024","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7501807,0.003284073,0.2204907,0.0020365706,0.00029422075,0.0010025991,0.0025398172,0.013479294,0.0066919965],"genre_scores_gemma":[0.7838063,0.00040980498,0.2089282,0.0001647446,0.000079633384,0.00026855897,0.0048521655,0.00048631005,0.0010042492],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.94567585,0.032830328,0.0055846083,0.004168293,0.010917434,0.0008235121],"domain_scores_gemma":[0.62685436,0.25355804,0.026148207,0.028875144,0.061459117,0.00310518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06020638,0.001616496,0.0011700172,0.0076394295,0.00086710474,0.0054017613,0.002703478,0.001692392,0.0009202421],"category_scores_gemma":[0.2800744,0.00063070253,0.0008727176,0.0032718629,0.0008592391,0.0052163536,0.0026109074,0.00149396,0.0009474859],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011745176,0.0009023637,0.15663259,0.00094624236,0.00041785723,0.00016938856,0.0031954523,0.092197426,0.008149317,0.0029595764,0.009857561,0.7233976],"study_design_scores_gemma":[0.00027900955,0.0021055567,0.075416334,0.00025928862,0.0003269765,0.00048134042,0.0017350329,0.8548175,0.0419348,0.0062435903,0.016108565,0.00029186014],"about_ca_topic_score_codex":0.0066641383,"about_ca_topic_score_gemma":0.007066231,"teacher_disagreement_score":0.06020638,"about_ca_system_score_codex":0.0025443716,"about_ca_system_score_gemma":0.003567552,"threshold_uncertainty_score":0.31840557},"labels":[],"label_agreement":null},{"id":"W4385550818","doi":"10.31219/osf.io/kn5f2","title":"Performance Analysis of Large Language Models for Medical Text Summarization","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Automatic summarization; Pace; Relevance (law); Computer science; Unified Medical Language System; Medical literature; Data science; Natural language processing; Medicine; Pathology","score_opus":0.03542364612472997,"score_gpt":0.30563387709382067,"score_spread":0.2702102309690907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385550818","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.66425717,0.024241067,0.22584352,0.0055229417,0.002097622,0.0010379867,0.02043073,0.050153725,0.0064151986],"genre_scores_gemma":[0.84552336,0.002136234,0.111956656,0.00064598897,0.0005290581,0.00045534776,0.035081167,0.0007683379,0.0029038903],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99525636,0.0026327036,0.00058031804,0.00082548714,0.00044077652,0.00026434727],"domain_scores_gemma":[0.9767104,0.01814943,0.00082453556,0.0013832585,0.002370274,0.00056207133],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0095321955,0.002494972,0.0016288582,0.0032808713,0.0007285341,0.0017116423,0.0016193498,0.0018938656,0.001841204],"category_scores_gemma":[0.030781424,0.0004812392,0.0015303922,0.0022952973,0.00047236454,0.0027515013,0.0010712664,0.0019400996,0.0017374197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0052258032,0.0012031476,0.012778225,0.0024837519,0.0019813736,0.0006549786,0.00062443153,0.47929758,0.022089377,0.0017735912,0.044551864,0.42733583],"study_design_scores_gemma":[0.000158286,0.00062228384,0.0023872545,0.00005144158,0.00026083487,0.000104587336,0.00016143224,0.98321784,0.009133682,0.0014184006,0.0024261125,0.00005786127],"about_ca_topic_score_codex":0.016069094,"about_ca_topic_score_gemma":0.015493302,"teacher_disagreement_score":0.016069094,"about_ca_system_score_codex":0.0017547206,"about_ca_system_score_gemma":0.0020756193,"threshold_uncertainty_score":0.0504117},"labels":[],"label_agreement":null},{"id":"W4385565351","doi":"10.18653/v1/2023.acl-long.99","title":"Precise Zero-Shot Dense Retrieval without Relevance Labels","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":233,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Relevance (law); Computer science; Similarity (geometry); Relevance feedback; Embedding; Vector space model; Artificial intelligence; Information retrieval; Encoder; Zero (linguistics); Natural language processing; Document retrieval; Vector space; Encoding (memory); Language model; Image retrieval; Pattern recognition (psychology); Mathematics; Image (mathematics)","score_opus":0.051737550002513226,"score_gpt":0.29283138815349513,"score_spread":0.2410938381509819,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385565351","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053525466,0.0010466998,0.9326018,0.0002709103,0.0000831902,0.0002569932,0.00036889294,0.0077665458,0.004079486],"genre_scores_gemma":[0.58717704,0.00059579685,0.39427882,0.00059075595,0.00013269579,0.0002461056,0.0020844454,0.000523384,0.014370839],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988398,0.00024583837,0.00007671615,0.0003577325,0.0003377027,0.00014229487],"domain_scores_gemma":[0.997454,0.00095349643,0.00016296943,0.0009844882,0.0003374449,0.000107598564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016217938,0.0008307018,0.0012867294,0.0008250775,0.000632689,0.0012648285,0.0022476376,0.0011222846,0.0043313005],"category_scores_gemma":[0.006049732,0.00054869126,0.0005497888,0.0008188074,0.0012786668,0.0051054345,0.0030874733,0.0013669534,0.002513336],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093395985,0.00058706006,0.0016821172,0.0006989691,0.00012852032,0.00026556148,0.0006928482,0.0389984,0.073282294,0.02010466,0.013724105,0.8489015],"study_design_scores_gemma":[0.00020344096,0.00094757864,0.0021815635,0.00007830985,0.00013076504,0.0012050655,0.0005656751,0.8439937,0.08150802,0.053129394,0.015918791,0.00013766147],"about_ca_topic_score_codex":0.0039436514,"about_ca_topic_score_gemma":0.0075307363,"teacher_disagreement_score":0.0043313005,"about_ca_system_score_codex":0.00077437604,"about_ca_system_score_gemma":0.0012435829,"threshold_uncertainty_score":0.014489591},"labels":[],"label_agreement":null},{"id":"W4385565506","doi":"10.18653/v1/2023.bea-1.62","title":"Empowering Conversational Agents using Semantic In-Context Learning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Task (project management); Computer science; Context (archaeology); Work (physics); Artificial intelligence; Data science; Natural language processing; Human–computer interaction; Knowledge management; Engineering","score_opus":0.061271369527666014,"score_gpt":0.3122385994401016,"score_spread":0.2509672299124356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385565506","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04836032,0.0010641192,0.9281862,0.0022609688,0.0002510851,0.00038985477,0.0003549553,0.00860515,0.010527334],"genre_scores_gemma":[0.61150616,0.0005812976,0.38020444,0.00073850923,0.00015061714,0.00040558458,0.0008165468,0.0003674061,0.0052294508],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99654466,0.0021078335,0.00014640475,0.00058685313,0.0004064798,0.00020779636],"domain_scores_gemma":[0.99505275,0.003177899,0.00022920578,0.00075119786,0.00046397728,0.00032510413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004009705,0.0012289317,0.0007976681,0.00075681566,0.0010499241,0.00277468,0.0017958003,0.0016901209,0.0032628481],"category_scores_gemma":[0.011926873,0.0005579702,0.00095399836,0.0004971514,0.0010019743,0.0078113936,0.004826333,0.0034944124,0.0019599434],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00097328634,0.0018839348,0.010864637,0.0012883552,0.00035165853,0.0009456339,0.008668944,0.10100733,0.043524016,0.06651105,0.019634563,0.7443466],"study_design_scores_gemma":[0.00006074423,0.00022947429,0.00066366146,0.000102877115,0.000114997354,0.00020349749,0.001625008,0.8617195,0.016597863,0.083423726,0.03517173,0.000086941924],"about_ca_topic_score_codex":0.0027444072,"about_ca_topic_score_gemma":0.005276804,"teacher_disagreement_score":0.004009705,"about_ca_system_score_codex":0.00064614724,"about_ca_system_score_gemma":0.0016005763,"threshold_uncertainty_score":0.021205604},"labels":[],"label_agreement":null},{"id":"W4385566909","doi":"10.18653/v1/2022.nlpcss-1.17","title":"An Analysis of Acknowledgments in NLP Conference Proceedings","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Army Research Office; National Key Research and Development Program of China; Engineering and Physical Sciences Research Council; Air Force Research Laboratory; Intelligence Advanced Research Projects Activity; Defense Advanced Research Projects Agency; Office of Naval Research; National Natural Science Foundation of China; Advanced Research Projects Agency; Impact Fund; National Institutes of Health; Nvidia; Atomic Energy of Canada Limited","keywords":"Computer science; Field (mathematics); Section (typography); Artificial intelligence; Natural language processing; Default; State (computer science); Incentive; Data science; Programming language; Mathematics","score_opus":0.031898717368553785,"score_gpt":0.27910495483539316,"score_spread":0.24720623746683937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385566909","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6907068,0.02867359,0.09596744,0.01449578,0.00293764,0.0007679599,0.07730651,0.0035079604,0.085636266],"genre_scores_gemma":[0.91135526,0.00489195,0.033604633,0.0005465354,0.0019522925,0.0007819681,0.03421687,0.0007685809,0.011881892],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9815196,0.005652134,0.0024427278,0.0021572886,0.007472021,0.0007562734],"domain_scores_gemma":[0.7653068,0.16577211,0.033318,0.008205277,0.024390353,0.003007486],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.011991381,0.00042112285,0.00048249893,0.016475923,0.0021481968,0.0047993697,0.0010648298,0.0008414761,0.0064107394],"category_scores_gemma":[0.1175614,0.00033287614,0.00051090156,0.024202283,0.0011362057,0.0056593125,0.0018922988,0.001330358,0.002253354],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010645327,0.0002829274,0.20196564,0.0057263854,0.0005000044,0.0023062602,0.03176214,0.009394509,0.012368367,0.08775867,0.14783382,0.49903673],"study_design_scores_gemma":[0.00009796209,0.00032533394,0.3004196,0.0021231319,0.00038826576,0.0033633793,0.020339163,0.025209965,0.008479833,0.05127487,0.58766085,0.00031757308],"about_ca_topic_score_codex":0.0023285954,"about_ca_topic_score_gemma":0.002975478,"teacher_disagreement_score":0.9880086,"about_ca_system_score_codex":0.0014979651,"about_ca_system_score_gemma":0.001662955,"threshold_uncertainty_score":0.063417256},"labels":[],"label_agreement":null},{"id":"W4385567355","doi":"10.18653/v1/2022.findings-emnlp.502","title":"Probing for Constituency Structure in Neural Language Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Deutsche Forschungsgemeinschaft","keywords":"Treebank; Computer science; Natural language processing; Focus (optics); Artificial intelligence; Syntactic structure; Syntax; Tree structure; Parsing; Data structure; Programming language","score_opus":0.023515779029376554,"score_gpt":0.2507793196022692,"score_spread":0.22726354057289264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385567355","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37166008,0.0005810472,0.622031,0.000705205,0.000062883635,0.000046035577,0.00029795244,0.0020208866,0.0025949734],"genre_scores_gemma":[0.9368664,0.00019566475,0.060154215,0.000146327,0.00003312342,0.000095205905,0.0007386128,0.0001799,0.001590593],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993395,0.00030257922,0.00002523903,0.00023151017,0.00005128004,0.000049894275],"domain_scores_gemma":[0.99714845,0.0022176048,0.00014584178,0.00032123696,0.000108830856,0.00005793031],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017180404,0.0007033273,0.0006138295,0.00046666522,0.0004209075,0.0010622233,0.0008790968,0.0009853326,0.0020242936],"category_scores_gemma":[0.00918738,0.0005700247,0.00076115824,0.00059407693,0.00068980485,0.0040237657,0.0010856459,0.002795945,0.0006314495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009357755,0.0002322872,0.013558157,0.00040945935,0.00034526066,0.00031223428,0.0013248378,0.49044377,0.06642645,0.05124125,0.004147502,0.37062302],"study_design_scores_gemma":[0.00000677926,0.000048145168,0.001076027,0.0000104392175,0.000022647433,0.000028747341,0.00004610613,0.9752333,0.003614354,0.019310668,0.0005924397,0.000010237451],"about_ca_topic_score_codex":0.002056199,"about_ca_topic_score_gemma":0.0042749476,"teacher_disagreement_score":0.002056199,"about_ca_system_score_codex":0.00069636304,"about_ca_system_score_gemma":0.00048910803,"threshold_uncertainty_score":0.009086013},"labels":[],"label_agreement":null},{"id":"W4385567372","doi":"10.18653/v1/2022.findings-emnlp.256","title":"RoChBert: Towards Robust BERT Fine-tuning for Chinese","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Key Research and Development Program of China; Beijing National Research Center For Information Science And Technology; National Natural Science Foundation of China","keywords":"Computer science; Robustness (evolution); Adversarial system; Artificial intelligence; Glyph (data visualization); Language model; Scratch; Fuse (electrical); Machine learning; Natural language processing; Visualization; Programming language","score_opus":0.0326339763209775,"score_gpt":0.25944590716233257,"score_spread":0.22681193084135506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385567372","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08662623,0.0015598806,0.8899896,0.00062895956,0.00026090775,0.0002069769,0.0007268633,0.01273498,0.0072655557],"genre_scores_gemma":[0.7976892,0.00064896204,0.18133222,0.0007253603,0.00012271348,0.00029264222,0.002493247,0.0008805324,0.015815109],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936455,0.00018535096,0.00002631974,0.00019932,0.00012871115,0.00009583902],"domain_scores_gemma":[0.99926454,0.00030535727,0.000063789426,0.00021434306,0.00010134355,0.000050667542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014523767,0.001627776,0.00092477445,0.0005967882,0.0005240261,0.00079148664,0.0015223215,0.0010545679,0.0028050053],"category_scores_gemma":[0.002906504,0.0005621671,0.0009623964,0.0004446914,0.0009671006,0.001949469,0.0018499786,0.0021036193,0.0016333252],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030114606,0.00014021927,0.0015242485,0.00017420325,0.00012231356,0.0002392822,0.00020698909,0.7216398,0.02263128,0.015249801,0.011383157,0.2263876],"study_design_scores_gemma":[0.000013234256,0.000051128034,0.00011515212,0.00000718977,0.000008236166,0.000033736444,0.000012973802,0.99003863,0.0041920426,0.003934028,0.0015810905,0.000012613194],"about_ca_topic_score_codex":0.008848082,"about_ca_topic_score_gemma":0.01214198,"teacher_disagreement_score":0.008848082,"about_ca_system_score_codex":0.0009038558,"about_ca_system_score_gemma":0.0012199228,"threshold_uncertainty_score":0.017593145},"labels":[],"label_agreement":null},{"id":"W4385567884","doi":"10.1145/3580305.3599411","title":"Learning to Relate to Previous Turns in Conversational Search","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Natural Science Foundation of China","keywords":"Computer science; Conversation; Information retrieval; Task (project management); Query expansion; Selection (genetic algorithm); Context (archaeology); Web search query; Query language; Search engine; Artificial intelligence","score_opus":0.040873855434147476,"score_gpt":0.2949330662295956,"score_spread":0.25405921079544813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385567884","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36242822,0.00983003,0.6117699,0.0014745663,0.0002065704,0.00051696436,0.0012461359,0.00500047,0.007527158],"genre_scores_gemma":[0.9170279,0.0009453468,0.074298784,0.0004775949,0.00028005653,0.00031493217,0.0021338693,0.00026797215,0.0042534275],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99613357,0.0018659004,0.00018942016,0.0011131329,0.0003496592,0.0003483387],"domain_scores_gemma":[0.9906545,0.0069704233,0.00061417895,0.0006925792,0.0007533149,0.0003149443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051568504,0.001755648,0.0015743895,0.0024553684,0.0012269941,0.0015968683,0.0022439428,0.0020438372,0.0016548464],"category_scores_gemma":[0.016234545,0.0008035532,0.0013777806,0.0018686854,0.00090297865,0.0042919274,0.0020755294,0.0019919777,0.001465415],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029295506,0.0011282861,0.044098604,0.0015218685,0.00061859365,0.0007122211,0.0050812094,0.124718346,0.030263584,0.006225096,0.018774746,0.763928],"study_design_scores_gemma":[0.0000984859,0.00036281405,0.007518449,0.00007597174,0.00023402928,0.00038955236,0.00065561314,0.96640736,0.0059911353,0.013366794,0.004802917,0.00009695975],"about_ca_topic_score_codex":0.012399272,"about_ca_topic_score_gemma":0.017427627,"teacher_disagreement_score":0.012399272,"about_ca_system_score_codex":0.0010446138,"about_ca_system_score_gemma":0.0015349026,"threshold_uncertainty_score":0.027272344},"labels":[],"label_agreement":null},{"id":"W4385569664","doi":"10.18653/v1/2023.acl-long.341","title":"Forgotten Knowledge: Examining the Citational Amnesia in NLP","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Citation; Diversity (politics); Computer science; Reading (process); Set (abstract data type); Artificial intelligence; Lexical diversity; Data science; Information retrieval; Library science; Linguistics; Sociology; Vocabulary","score_opus":0.07681843267050391,"score_gpt":0.2994837716378186,"score_spread":0.22266533896731466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385569664","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96451557,0.0075188135,0.009263199,0.0023493357,0.000106134336,0.000070811104,0.0033350964,0.00018435637,0.012656631],"genre_scores_gemma":[0.9895767,0.00203885,0.004119315,0.00017502255,0.00020562603,0.000084280495,0.0019186643,0.00008070128,0.0018008967],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9940745,0.0014815782,0.0010148786,0.0010076244,0.0020842932,0.00033705842],"domain_scores_gemma":[0.8248467,0.12024886,0.034195088,0.0076506585,0.010960224,0.002098569],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.011252436,0.0002614816,0.0004875208,0.017813189,0.0015840002,0.004422269,0.0011119944,0.00092771143,0.0029191787],"category_scores_gemma":[0.11561747,0.00023468384,0.0005263792,0.041509848,0.0019610585,0.0070004836,0.0031919945,0.0012429027,0.00073088636],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017837889,0.00009252771,0.8540768,0.0009715654,0.00032249762,0.00083499757,0.020312395,0.0017731688,0.0012453726,0.011312357,0.0048799687,0.10400013],"study_design_scores_gemma":[0.000022312988,0.00010621385,0.91702425,0.0006439693,0.00027131144,0.0014278842,0.011234307,0.008574446,0.0018852764,0.028543193,0.030186558,0.00008029532],"about_ca_topic_score_codex":0.011437236,"about_ca_topic_score_gemma":0.011796054,"teacher_disagreement_score":0.98874754,"about_ca_system_score_codex":0.0018220519,"about_ca_system_score_gemma":0.0016488456,"threshold_uncertainty_score":0.059509277},"labels":[],"label_agreement":null},{"id":"W4385569686","doi":"10.18653/v1/2023.acl-long.274","title":"ConvGQR: Generative Query Reformulation for Conversational Search","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"China Scholarship Council; Tsinghua University","keywords":"Computer science; Query expansion; Conversation; Query language; Rewriting; Web search query; Information retrieval; Web query classification; Query optimization; RDF query language; Sargable; Generative grammar; Task (project management); Artificial intelligence; Search engine; Programming language; Linguistics","score_opus":0.08974196501350201,"score_gpt":0.32235838613890366,"score_spread":0.23261642112540165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385569686","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016812952,0.0023853492,0.95483255,0.00053495285,0.00011290336,0.00033745315,0.001588049,0.020588158,0.0028075757],"genre_scores_gemma":[0.3698982,0.0013903044,0.60477144,0.0015466989,0.00024879415,0.0007224755,0.009877419,0.0024819232,0.00906275],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983398,0.0006194072,0.00009738562,0.0004911748,0.00030894618,0.00014334594],"domain_scores_gemma":[0.99841154,0.0008740931,0.00008001952,0.0003458609,0.00022015456,0.00006829891],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018421138,0.0021725758,0.001888786,0.0013769247,0.00057268084,0.0011314787,0.003069192,0.0020581842,0.005994147],"category_scores_gemma":[0.005778874,0.0007359108,0.0022654901,0.0012390917,0.0010334571,0.002957377,0.0024060716,0.0024769478,0.0038221555],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072591944,0.00038537412,0.001573398,0.0013449294,0.00026133674,0.0006729051,0.0013456203,0.17590944,0.03783297,0.019030193,0.064351656,0.69656616],"study_design_scores_gemma":[0.00007263587,0.00015324107,0.00024009842,0.00002725477,0.00005906319,0.0002586075,0.00013175921,0.96479976,0.008307786,0.017399581,0.008506636,0.000043629945],"about_ca_topic_score_codex":0.012112361,"about_ca_topic_score_gemma":0.014937905,"teacher_disagreement_score":0.012112361,"about_ca_system_score_codex":0.0014088858,"about_ca_system_score_gemma":0.001500612,"threshold_uncertainty_score":0.024083674},"labels":[],"label_agreement":null},{"id":"W4385569777","doi":"10.18653/v1/2023.woah-1.14","title":"Concept-Based Explanations to Test for False Causal Relationships Learned by Abusive Language Classifiers","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Computer science; Artificial intelligence; Set (abstract data type); Feature (linguistics); Machine learning; Focus (optics); Natural language processing; Linguistics","score_opus":0.07954067804392582,"score_gpt":0.3169069890658719,"score_spread":0.23736631102194605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385569777","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7951541,0.0019640678,0.1893397,0.00226478,0.00038946513,0.00073975726,0.0027406479,0.001672132,0.005735273],"genre_scores_gemma":[0.96436566,0.00010883428,0.03194629,0.00032938187,0.00009163341,0.00030837275,0.002218993,0.00010607838,0.000524831],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9830442,0.008579748,0.0014476045,0.002830446,0.0033239382,0.00077405694],"domain_scores_gemma":[0.6735955,0.28869584,0.01263827,0.014396604,0.009008816,0.0016649631],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.033811323,0.0018074294,0.0012013876,0.003025094,0.0011214868,0.0023992574,0.0019235407,0.0033306251,0.003611232],"category_scores_gemma":[0.1607063,0.00037694265,0.0012611005,0.001721695,0.002736553,0.0061089583,0.0031500207,0.0058917366,0.0006984463],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005156829,0.0018415687,0.5194976,0.0018469047,0.0023780472,0.0011683562,0.0052162027,0.09796228,0.0116573125,0.028796995,0.01904866,0.3054292],"study_design_scores_gemma":[0.00043949217,0.002298187,0.115057796,0.00053071656,0.0009023936,0.0016136989,0.0031038146,0.75362635,0.029809736,0.07925034,0.0130827725,0.00028474448],"about_ca_topic_score_codex":0.0013902904,"about_ca_topic_score_gemma":0.0016662832,"teacher_disagreement_score":0.033811323,"about_ca_system_score_codex":0.0010561524,"about_ca_system_score_gemma":0.0012355332,"threshold_uncertainty_score":0.17881346},"labels":[],"label_agreement":null},{"id":"W4385569780","doi":"10.18653/v1/2023.acl-long.307","title":"Evaluating Open-Domain Question Answering in the Era of Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":99,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Matching (statistics); Benchmark (surveying); Question answering; Open domain; Natural language processing; Domain (mathematical analysis); Artificial intelligence; Language model; Information retrieval; Machine learning; Statistics","score_opus":0.07837739210868339,"score_gpt":0.3860292843682005,"score_spread":0.3076518922595171,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385569780","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39410588,0.00784396,0.5459175,0.0043111984,0.0006794541,0.00076207565,0.004411997,0.029099695,0.012868299],"genre_scores_gemma":[0.8275643,0.0005027548,0.1586239,0.0011389261,0.00017895461,0.00029180667,0.008642368,0.0007809546,0.0022759924],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9834009,0.010382399,0.000706792,0.0024380297,0.0026915735,0.000380287],"domain_scores_gemma":[0.95550853,0.03268451,0.0011374955,0.00547475,0.0040704273,0.0011243179],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015959272,0.0019063414,0.0014339555,0.0018669147,0.0009443567,0.003375608,0.003042781,0.0030932366,0.0033565592],"category_scores_gemma":[0.06972561,0.0006116046,0.0013465335,0.0011127595,0.0014576571,0.006397056,0.0040117507,0.0038257518,0.0022404152],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016416097,0.0014579484,0.024136132,0.003061755,0.0009552458,0.00058526476,0.0029908097,0.3154905,0.025684802,0.024130873,0.041915324,0.55794966],"study_design_scores_gemma":[0.00009332826,0.00039736996,0.0029634312,0.00011991565,0.00009079894,0.00018256273,0.00047470437,0.95394546,0.009180645,0.022211013,0.010284802,0.000056056553],"about_ca_topic_score_codex":0.0074605304,"about_ca_topic_score_gemma":0.010616519,"teacher_disagreement_score":0.015959272,"about_ca_system_score_codex":0.0022780641,"about_ca_system_score_gemma":0.0019031633,"threshold_uncertainty_score":0.08440173},"labels":[],"label_agreement":null},{"id":"W4385569905","doi":"10.18653/v1/2023.woah-1.16","title":"Conversation Derailment Forecasting with Graph Convolutional Networks","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Conversation; Computer science; Derailment; Benchmark (surveying); Convolutional neural network; Moderation; Graph; Artificial intelligence; Machine learning; Psychology; Theoretical computer science; Communication","score_opus":0.04109949231874602,"score_gpt":0.21991786078820816,"score_spread":0.17881836846946214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385569905","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60119826,0.005836024,0.36852044,0.002592631,0.0005766314,0.00021034868,0.0065585845,0.00635022,0.008156894],"genre_scores_gemma":[0.9718137,0.00041647113,0.020528514,0.00014400299,0.00013495999,0.00006947026,0.0038825155,0.000063686784,0.0029467954],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.999617,0.00010247517,0.000017690636,0.00014395962,0.00004985937,0.00006910131],"domain_scores_gemma":[0.99884593,0.0006700448,0.00012485716,0.000094895855,0.00019717128,0.00006697343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073312945,0.0012121265,0.00053151726,0.0011853081,0.0004264533,0.0006742117,0.0010786672,0.00090989645,0.0011542548],"category_scores_gemma":[0.003246272,0.0003435934,0.00065179076,0.00092325313,0.00030069318,0.0010940148,0.00068379554,0.0014662377,0.00059455214],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007549338,0.00035157538,0.01980539,0.00017113774,0.00025966804,0.0002290788,0.00032208956,0.72127694,0.005709019,0.003907265,0.01637157,0.23084144],"study_design_scores_gemma":[0.0000038109592,0.000009485671,0.00085487514,0.0000038646135,0.000011454328,0.000007155423,0.000014581075,0.99715173,0.00036051052,0.00116586,0.00041233576,0.000004285242],"about_ca_topic_score_codex":0.035810426,"about_ca_topic_score_gemma":0.045775283,"teacher_disagreement_score":0.035810426,"about_ca_system_score_codex":0.0013593232,"about_ca_system_score_gemma":0.0008195887,"threshold_uncertainty_score":0.07120395},"labels":[],"label_agreement":null},{"id":"W4385570008","doi":"10.18653/v1/2023.acl-long.770","title":"Optimal Transport for Unsupervised Hallucination Detection in Neural Machine Translation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fundação para a Ciência e a Tecnologia; HORIZON EUROPE Framework Programme; Grand Équipement National De Calcul Intensif","keywords":"Computer science; Machine translation; Detector; Intuition; Artificial intelligence; Sentence; Autoencoder; De facto; Translation (biology); Similarity (geometry); Machine learning; Artificial neural network; Natural language processing","score_opus":0.0426722949187765,"score_gpt":0.26741477469606145,"score_spread":0.22474247977728495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570008","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02007106,0.00032945545,0.9758354,0.00036675786,0.00010276712,0.0000421128,0.0002166968,0.002055816,0.0009798988],"genre_scores_gemma":[0.57175,0.00048607728,0.41111758,0.00026263171,0.0003235386,0.00020248945,0.0019851073,0.0015769525,0.012295577],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99926025,0.00026385082,0.000046368838,0.0001945575,0.00011592105,0.00011907498],"domain_scores_gemma":[0.9965758,0.0021306688,0.00019029458,0.0003537311,0.0005920187,0.0001575249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017051493,0.0010856036,0.001661246,0.001093245,0.0011453638,0.001667372,0.001430865,0.002090763,0.005330185],"category_scores_gemma":[0.008973391,0.00073559704,0.0011865259,0.0012208393,0.001052179,0.0023923228,0.0021531435,0.0020700465,0.0020625552],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012196393,0.00021128039,0.002166036,0.00035033558,0.00013673394,0.00033512423,0.00039694304,0.41745618,0.017601853,0.034727827,0.011880804,0.51351726],"study_design_scores_gemma":[0.000011893062,0.00003058377,0.00015042552,0.000010639878,0.000007962177,0.000033227352,0.000025363375,0.9834704,0.0026680718,0.012983305,0.00059929397,0.000008801753],"about_ca_topic_score_codex":0.0077592144,"about_ca_topic_score_gemma":0.0064399964,"teacher_disagreement_score":0.0077592144,"about_ca_system_score_codex":0.0011436472,"about_ca_system_score_gemma":0.0015562034,"threshold_uncertainty_score":0.017831206},"labels":[],"label_agreement":null},{"id":"W4385570025","doi":"10.18653/v1/2023.acl-long.385","title":"Few-shot In-context Learning on Knowledge Base Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Knowledge base; Question answering; Executable; Schema (genetic algorithms); Artificial intelligence; Context (archaeology); Natural language; Natural language processing; Entity linking; Natural language understanding; Matching (statistics); Baseline (sea); Information retrieval; Programming language","score_opus":0.050013981675477326,"score_gpt":0.30104167190427256,"score_spread":0.25102769022879523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570025","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05328493,0.0050566536,0.91727275,0.0014881798,0.0002969436,0.00038910596,0.0015666591,0.01684937,0.00379543],"genre_scores_gemma":[0.59427977,0.0011375678,0.38740543,0.001883056,0.00040246174,0.00047336094,0.008382041,0.0007100949,0.0053262],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99795437,0.0007642198,0.00010502327,0.0008019946,0.0002222075,0.00015212341],"domain_scores_gemma":[0.9953797,0.0032299592,0.00012632927,0.000632778,0.00042518435,0.0002060864],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026589632,0.0018732519,0.0016697516,0.0018345395,0.000984301,0.0015423569,0.0038355724,0.002998902,0.0049440903],"category_scores_gemma":[0.009699998,0.0008786952,0.0016285442,0.0013758779,0.0011379828,0.0050272667,0.0027260997,0.0037308165,0.0020451965],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006556789,0.00076664967,0.0036820457,0.001176715,0.0003834473,0.00052692567,0.0010486856,0.23943849,0.012404206,0.011441641,0.02936997,0.6991055],"study_design_scores_gemma":[0.00004806577,0.00011841133,0.00063645013,0.000051619198,0.000061359206,0.00013715614,0.0001417436,0.9683071,0.0036082505,0.021834435,0.005027392,0.000028041437],"about_ca_topic_score_codex":0.012955981,"about_ca_topic_score_gemma":0.018231202,"teacher_disagreement_score":0.012955981,"about_ca_system_score_codex":0.001460379,"about_ca_system_score_gemma":0.001176246,"threshold_uncertainty_score":0.025761187},"labels":[],"label_agreement":null},{"id":"W4385570159","doi":"10.18653/v1/2023.findings-acl.714","title":"A Study on Knowledge Distillation from Weak Teacher for Scaling Up Pre-trained Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Initialization; Computer science; Artificial intelligence; Weighting; Distillation; Machine learning; Language model; Domain (mathematical analysis); Natural language processing; Mathematics","score_opus":0.0692850980979679,"score_gpt":0.3390118085480475,"score_spread":0.2697267104500796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570159","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31106135,0.0022875725,0.6740603,0.0012003691,0.00019177113,0.00034671064,0.00018061059,0.004428124,0.006243194],"genre_scores_gemma":[0.66232723,0.00064695807,0.33214873,0.00029605723,0.00007191708,0.00016927299,0.0005389426,0.00050638244,0.0032945292],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99900997,0.0003887283,0.00008550315,0.00023645443,0.00019892916,0.00008044573],"domain_scores_gemma":[0.9937884,0.004112816,0.00020270147,0.0008883737,0.00081646873,0.00019115147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003323753,0.0011508124,0.00086063286,0.00060322654,0.00047924492,0.001248088,0.0015745771,0.0010961953,0.0021118592],"category_scores_gemma":[0.018269563,0.0006309454,0.0007337781,0.0008040514,0.0006880295,0.0041806474,0.0016552891,0.0026713973,0.0005911286],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007926587,0.00065705006,0.005227523,0.0005476104,0.00017022072,0.00022598042,0.0005979614,0.39537242,0.034101605,0.009382524,0.0030553539,0.5498691],"study_design_scores_gemma":[0.000052401138,0.00032658517,0.0005873047,0.000022694725,0.000037283808,0.00006611811,0.000066862085,0.9795306,0.015522541,0.001621919,0.0021524068,0.000013324236],"about_ca_topic_score_codex":0.006189129,"about_ca_topic_score_gemma":0.006056724,"teacher_disagreement_score":0.006189129,"about_ca_system_score_codex":0.0009882589,"about_ca_system_score_gemma":0.0014514136,"threshold_uncertainty_score":0.017577887},"labels":[],"label_agreement":null},{"id":"W4385570165","doi":"10.18653/v1/2023.findings-acl.206","title":"Two Examples are Better than One: Context Regularization for Gradient-based Prompt Tuning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Regularization (linguistics); Context (archaeology); Machine learning; Artificial intelligence; Language model","score_opus":0.06752832518457182,"score_gpt":0.2678553627134493,"score_spread":0.20032703752887748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570165","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0718805,0.0020180806,0.90062636,0.0015580767,0.0004281762,0.00015584832,0.00019448648,0.018904686,0.004233788],"genre_scores_gemma":[0.55471975,0.00031745527,0.43850648,0.0010505493,0.0001353741,0.00019654182,0.0004840453,0.0010269397,0.0035628162],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983524,0.0006052065,0.000078920966,0.0005889613,0.00024919692,0.00012538358],"domain_scores_gemma":[0.9977543,0.0010769612,0.00011621311,0.0005743411,0.00029150012,0.00018669375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002035459,0.0016340603,0.0010847512,0.00046114696,0.0007260958,0.0012069471,0.0017760559,0.0023163215,0.0028052283],"category_scores_gemma":[0.010384556,0.0005533279,0.0007388791,0.0003674699,0.0011825513,0.0030970646,0.0021932945,0.0045593367,0.0012644732],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019808165,0.0010262878,0.0051141446,0.0006154607,0.0002138108,0.00037703817,0.00095669116,0.22667153,0.056009322,0.02569846,0.028078917,0.6532576],"study_design_scores_gemma":[0.00016351296,0.0003548545,0.00073872384,0.000049094935,0.000041975425,0.00017149858,0.000083884064,0.9590073,0.014773737,0.018606199,0.0059368634,0.00007227005],"about_ca_topic_score_codex":0.0023591376,"about_ca_topic_score_gemma":0.004834129,"teacher_disagreement_score":0.0028052283,"about_ca_system_score_codex":0.0006179099,"about_ca_system_score_gemma":0.0011856485,"threshold_uncertainty_score":0.010764658},"labels":[],"label_agreement":null},{"id":"W4385570189","doi":"10.18653/v1/2023.findings-acl.139","title":"This prompt is measuring &lt;mask&gt;: evaluating bias evaluation in language models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"University of Edinburgh; UK Research and Innovation","keywords":"Computer science; Scope (computer science); Measure (data warehouse); Taxonomy (biology); Language model; Field (mathematics); Gender bias; Natural language processing; Artificial intelligence; Data science; Psychology; Data mining; Social psychology; Mathematics","score_opus":0.29236968409882247,"score_gpt":0.36457375483565485,"score_spread":0.07220407073683238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570189","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26194587,0.0022843936,0.66442484,0.013305487,0.0012441027,0.004259778,0.0048772646,0.00598043,0.041677874],"genre_scores_gemma":[0.67481554,0.000438972,0.31358057,0.0023904317,0.00019426056,0.0032271212,0.0015863563,0.00085659645,0.0029101125],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.8748277,0.09473871,0.0067001325,0.005449112,0.01705811,0.0012262146],"domain_scores_gemma":[0.44933268,0.43800986,0.024357622,0.05661663,0.028531823,0.0031513507],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.094175324,0.0014523662,0.0016548345,0.0029745318,0.0022686115,0.008408043,0.0021386896,0.0045848517,0.008770748],"category_scores_gemma":[0.49973717,0.0007114463,0.0015432854,0.0030466672,0.005853253,0.012504363,0.0067328806,0.003948894,0.0021354589],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043547316,0.0010094082,0.08821746,0.004132224,0.000644497,0.00027701395,0.02325848,0.021046255,0.010394712,0.17047602,0.030922418,0.6452668],"study_design_scores_gemma":[0.0013103719,0.00512911,0.05239546,0.0039328537,0.0011294485,0.00086554314,0.012092738,0.14779077,0.047485247,0.6290632,0.09788457,0.00092064904],"about_ca_topic_score_codex":0.0035283726,"about_ca_topic_score_gemma":0.004065266,"teacher_disagreement_score":0.90582466,"about_ca_system_score_codex":0.0035812347,"about_ca_system_score_gemma":0.005732103,"threshold_uncertainty_score":0.49805266},"labels":[],"label_agreement":null},{"id":"W4385570227","doi":"10.18653/v1/2023.semeval-1.283","title":"Team TheSyllogist at SemEval-2023 Task 3: Language-Agnostic Framing Detection in Multi-Lingual Online News: A Zero-Shot Transfer Approach","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Framing (construction); Natural language processing; Artificial intelligence; SemEval; German; Machine translation; Georgian; Task (project management); Transfer of learning; Linguistics","score_opus":0.0601267243109582,"score_gpt":0.300719584713102,"score_spread":0.2405928604021438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570227","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25044033,0.0013272323,0.6245869,0.0018343775,0.0013266761,0.001040373,0.0062058405,0.079707794,0.033530474],"genre_scores_gemma":[0.6666607,0.00029379176,0.28312454,0.0008393368,0.00043683703,0.00077562104,0.012296474,0.0024361338,0.03313647],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9980294,0.00048327885,0.00007418152,0.0009301489,0.00026385023,0.00021906456],"domain_scores_gemma":[0.99768555,0.0007999302,0.00011332285,0.00059632637,0.00059698103,0.00020790276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024758254,0.0016167631,0.0011387266,0.0016812186,0.0011402387,0.0022472441,0.0015026625,0.0022146124,0.0069573917],"category_scores_gemma":[0.005992153,0.000503673,0.00092896394,0.00060369726,0.00054313306,0.0023954269,0.0025452564,0.0021814955,0.0077347113],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013665811,0.00091650145,0.009933091,0.00048421722,0.0002412329,0.0011983417,0.0025440997,0.009651944,0.12835498,0.0037010799,0.07812497,0.763483],"study_design_scores_gemma":[0.0002514183,0.00081537326,0.020430682,0.00013502485,0.0002700663,0.0017058448,0.0022592263,0.65912765,0.21246554,0.0129540935,0.089316815,0.00026831418],"about_ca_topic_score_codex":0.00429251,"about_ca_topic_score_gemma":0.005178089,"teacher_disagreement_score":0.0069573917,"about_ca_system_score_codex":0.0006871252,"about_ca_system_score_gemma":0.00095891283,"threshold_uncertainty_score":0.02327478},"labels":[],"label_agreement":null},{"id":"W4385570249","doi":"10.18653/v1/2023.findings-acl.177","title":"What Knowledge Is Needed? Towards Explainable Memory for kNN-MT Domain Adaptation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Novelis (Canada)","funders":"National Natural Science Foundation of China; National Science Foundation","keywords":"Correctness; Computer science; Interpretability; Security token; Artificial intelligence; Adaptation (eye); Pruning; Domain (mathematical analysis); Domain adaptation; Machine translation; Natural language processing; Machine learning; Theoretical computer science; Algorithm","score_opus":0.06855755797737949,"score_gpt":0.30038393913101663,"score_spread":0.23182638115363713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.066382274,0.00042233153,0.92659086,0.0017469232,0.000038831102,0.00010102672,0.00036315568,0.0025394373,0.0018151014],"genre_scores_gemma":[0.696974,0.00037824636,0.29902023,0.00035690045,0.000083945466,0.00020524701,0.0012401658,0.0005063593,0.0012349269],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99741757,0.001201594,0.00021918576,0.00062714587,0.00036049334,0.00017398802],"domain_scores_gemma":[0.9806467,0.01078488,0.0009862133,0.0059507336,0.0013033656,0.00032808303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004070393,0.0009433455,0.0012921257,0.001270219,0.0008562453,0.0028591852,0.0025529684,0.0014182613,0.0030121885],"category_scores_gemma":[0.03460214,0.0007938981,0.000929613,0.001221047,0.0019152529,0.009432675,0.003845268,0.002843852,0.0010458401],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00100425,0.0002965688,0.014740393,0.00086142024,0.00033840613,0.0012433896,0.0068231253,0.18074664,0.027764887,0.10205778,0.00674943,0.6573737],"study_design_scores_gemma":[0.000043580498,0.00008651326,0.0013871824,0.00009692686,0.000115121824,0.00026650503,0.0008988654,0.80423284,0.013443803,0.17242214,0.0069657443,0.00004073438],"about_ca_topic_score_codex":0.0042230086,"about_ca_topic_score_gemma":0.0053355647,"teacher_disagreement_score":0.0042230086,"about_ca_system_score_codex":0.0013578946,"about_ca_system_score_gemma":0.0019194154,"threshold_uncertainty_score":0.021526515},"labels":[],"label_agreement":null},{"id":"W4385570255","doi":"10.18653/v1/2023.acl-long.605","title":"f-Divergence Minimization for Sequence-Level Knowledge Distillation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; University of Alberta","funders":"Alliance de recherche numérique du Canada; DeepMind; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Divergence (linguistics); Distillation; Sequence (biology); Computer science; Minification; Function (biology); Process (computing); Word (group theory); Decomposition; Artificial intelligence; Mathematics; Programming language","score_opus":0.18341739238043794,"score_gpt":0.33519244037428825,"score_spread":0.1517750479938503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570255","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009950875,0.00048005002,0.9872573,0.00037512725,0.000038060083,0.000046802692,0.00015149156,0.0005691671,0.0011311853],"genre_scores_gemma":[0.42050204,0.0008292997,0.5678107,0.0007892955,0.00021247196,0.00037378023,0.0018968211,0.00054032775,0.0070452434],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981388,0.00063049805,0.0001314458,0.0004671051,0.00046058637,0.0001715564],"domain_scores_gemma":[0.9960854,0.0026192344,0.00019524843,0.0005058813,0.00042824997,0.00016594653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038635505,0.0012698572,0.0018269592,0.0013659949,0.0007869498,0.0015951488,0.0025138513,0.0023925856,0.0029712175],"category_scores_gemma":[0.01109343,0.00045142166,0.0010859345,0.0017601363,0.0019207753,0.00405794,0.002708633,0.0036607457,0.0011918594],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019205586,0.00022173725,0.0011373265,0.00029601163,0.000092515016,0.0001225931,0.00019263147,0.6269487,0.0035777893,0.06713248,0.008430584,0.2916556],"study_design_scores_gemma":[0.000012654176,0.000036931913,0.000116419804,0.000013654779,0.0000070698256,0.000036212543,0.000018793899,0.96026903,0.0013127917,0.036752842,0.0014104387,0.000013136758],"about_ca_topic_score_codex":0.0040026424,"about_ca_topic_score_gemma":0.005051699,"teacher_disagreement_score":0.0040026424,"about_ca_system_score_codex":0.0020574722,"about_ca_system_score_gemma":0.0025992945,"threshold_uncertainty_score":0.020432591},"labels":[],"label_agreement":null},{"id":"W4385570307","doi":"10.18653/v1/2023.acl-long.488","title":"DuNST: Dual Noisy Self Training for Semi-Supervised Controllable Text Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Controllability; Text generation; Fluency; Generalization; Artificial intelligence; Language model; Construct (python library); Process (computing); Exploit; Dual (grammatical number); Space (punctuation); Boundary (topology); Natural language processing; Machine learning; Mathematics","score_opus":0.07735932174670801,"score_gpt":0.27252835106490225,"score_spread":0.19516902931819424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570307","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030199464,0.0004186537,0.96011245,0.0002517605,0.00009714899,0.000119869794,0.0002198832,0.007058704,0.0015221188],"genre_scores_gemma":[0.61228764,0.00022972861,0.37527484,0.0007188923,0.00017540366,0.0005634192,0.0021195353,0.0012877843,0.0073427535],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989096,0.00040480087,0.000053738702,0.00038491783,0.0001633415,0.00008363075],"domain_scores_gemma":[0.9966691,0.0020601477,0.00018930632,0.00061118446,0.0003197231,0.00015056272],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018528547,0.0011600013,0.0009445426,0.0006455139,0.00053422316,0.0008226693,0.0024202631,0.0014944595,0.0026106988],"category_scores_gemma":[0.005469186,0.0006005939,0.0008587754,0.00050542573,0.0011801982,0.0022246577,0.0021990414,0.0021828867,0.0012659203],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050257245,0.0004659813,0.0025851682,0.000335624,0.00012392004,0.00025584208,0.00047366193,0.45848605,0.024014533,0.01268468,0.011147585,0.48892447],"study_design_scores_gemma":[0.00001909622,0.000051767667,0.0000910948,0.0000066787666,0.0000054532447,0.000024119221,0.000013144106,0.991852,0.003113497,0.004003194,0.00081254117,0.0000074477184],"about_ca_topic_score_codex":0.002013832,"about_ca_topic_score_gemma":0.004498079,"teacher_disagreement_score":0.0026106988,"about_ca_system_score_codex":0.0007261287,"about_ca_system_score_gemma":0.0008983317,"threshold_uncertainty_score":0.009798944},"labels":[],"label_agreement":null},{"id":"W4385570310","doi":"10.18653/v1/2023.acl-short.152","title":"A Simple and Effective Framework for Strict Zero-Shot Hierarchical Classification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Benchmark (surveying); Computer science; Task (project management); Contradiction; Zero (linguistics); Process (computing); Artificial intelligence; Machine learning; Shot (pellet); Simple (philosophy); Programming language; Geography; Engineering","score_opus":0.05952995306102866,"score_gpt":0.32867503133889736,"score_spread":0.2691450782778687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570310","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006252017,0.00060726097,0.9836559,0.00044593494,0.00014874262,0.00020511237,0.0005279188,0.0068189343,0.0013380607],"genre_scores_gemma":[0.2471038,0.0004720486,0.7356779,0.0011047502,0.0005640226,0.0007063382,0.0061941445,0.0011302079,0.00704679],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962094,0.0011876572,0.00026291556,0.001307222,0.00073540007,0.00029749313],"domain_scores_gemma":[0.9952427,0.0018946015,0.0003138907,0.0014043945,0.0008300695,0.00031432565],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045487364,0.002160523,0.0020691946,0.0029107083,0.002029338,0.0027557977,0.0052988906,0.0029096084,0.007032503],"category_scores_gemma":[0.014185436,0.0010000212,0.0018891133,0.0022553522,0.0015275052,0.0059983963,0.004779827,0.005975462,0.006483866],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005425844,0.00066802005,0.003822604,0.0004894254,0.00027199957,0.00029841758,0.00061765197,0.055754073,0.016869402,0.056555755,0.041901764,0.8222083],"study_design_scores_gemma":[0.000045003937,0.00009715826,0.0005303025,0.00004945092,0.000044138615,0.00014349601,0.00010224924,0.9118033,0.0042742915,0.07648399,0.0063743936,0.000052175317],"about_ca_topic_score_codex":0.0077564646,"about_ca_topic_score_gemma":0.016539412,"teacher_disagreement_score":0.0077564646,"about_ca_system_score_codex":0.00170317,"about_ca_system_score_gemma":0.0035653168,"threshold_uncertainty_score":0.024056315},"labels":[],"label_agreement":null},{"id":"W4385570349","doi":"10.18653/v1/2023.findings-acl.215","title":"Varta: A Large-Scale Headline-Generation Dataset for Indic Languages","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Canadian Institute for Advanced Research; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Headline; Computer science; Natural language processing; Artificial intelligence; Scale (ratio); Variety (cybernetics); Information retrieval; Data science; Linguistics; Geography","score_opus":0.05068896009991682,"score_gpt":0.3264950599578167,"score_spread":0.2758060998578999,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570349","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03430933,0.0014803006,0.0061093275,0.0006296998,0.00047658393,0.00042852302,0.929171,0.015301471,0.012093816],"genre_scores_gemma":[0.014009036,0.00021027395,0.009002594,0.00014533532,0.00007492881,0.00029268104,0.97351485,0.00048661232,0.0022637409],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99889636,0.00026716813,0.00015038026,0.000287462,0.00028406232,0.000114427494],"domain_scores_gemma":[0.9966395,0.001120751,0.0003174171,0.00066186645,0.00093333906,0.00032712313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010025869,0.001699349,0.00070606644,0.0050272853,0.0013640826,0.0014119503,0.0015461029,0.001619874,0.014345725],"category_scores_gemma":[0.00591154,0.00042416225,0.00095598557,0.0044741216,0.00052766595,0.0019479424,0.0015103424,0.001659681,0.01842492],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032836472,0.00029876342,0.006040018,0.0016231339,0.00009930126,0.0006334532,0.00065656693,0.0017390201,0.0048107556,0.0012808504,0.9374108,0.04507912],"study_design_scores_gemma":[0.00056910655,0.00022556014,0.024497695,0.00031269251,0.00015411907,0.0012988548,0.0016157296,0.019519832,0.012123428,0.0022367137,0.93726844,0.0001777951],"about_ca_topic_score_codex":0.01222862,"about_ca_topic_score_gemma":0.03175827,"teacher_disagreement_score":0.014345725,"about_ca_system_score_codex":0.0008624648,"about_ca_system_score_gemma":0.0017280978,"threshold_uncertainty_score":0.047991276},"labels":[],"label_agreement":null},{"id":"W4385570357","doi":"10.18653/v1/2023.findings-acl.533","title":"Fixed Input Parameterization for Efficient Prompting","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea","keywords":"Computer science; Inference; Task (project management); Parsing; Conversation; Fixed point; Overhead (engineering); Artificial intelligence; Property (philosophy); Theoretical computer science; Natural language processing; Programming language; Mathematics","score_opus":0.0487086013721159,"score_gpt":0.2818651687223084,"score_spread":0.23315656735019247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570357","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008183897,0.00013066111,0.9825362,0.00028387597,0.00006831019,0.000104611354,0.00017765012,0.0072831344,0.0012316555],"genre_scores_gemma":[0.36828816,0.00019257706,0.6228871,0.0005883188,0.00014014488,0.00073140167,0.0010847125,0.002013999,0.0040735938],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9963117,0.001880083,0.0002111957,0.00092213816,0.00041380964,0.0002610718],"domain_scores_gemma":[0.9884702,0.007629161,0.00032969273,0.0024385613,0.00085062487,0.0002817549],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037686364,0.0017818714,0.0014467996,0.00055401,0.000879659,0.0017312331,0.002973237,0.0023436954,0.015187373],"category_scores_gemma":[0.031576518,0.0009077313,0.0009098597,0.00074398465,0.0017871344,0.0058326377,0.004221331,0.004835501,0.005216997],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020194685,0.00047408883,0.002822375,0.0009975989,0.000095808544,0.0007825626,0.002918364,0.17862989,0.054737177,0.112673275,0.02255262,0.6212967],"study_design_scores_gemma":[0.00013743424,0.00020169569,0.0003738211,0.00007958113,0.000038184462,0.00018309586,0.00033726206,0.819307,0.023969073,0.13887365,0.016432462,0.000066723616],"about_ca_topic_score_codex":0.0013314037,"about_ca_topic_score_gemma":0.0019820868,"teacher_disagreement_score":0.015187373,"about_ca_system_score_codex":0.0011145938,"about_ca_system_score_gemma":0.0018105428,"threshold_uncertainty_score":0.05080688},"labels":[],"label_agreement":null},{"id":"W4385570371","doi":"10.18653/v1/2023.findings-acl.29","title":"A Systematic Study and Comprehensive Evaluation of ChatGPT on Benchmark Datasets","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":78,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Bank of Canada; York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Benchmark (surveying); Computer science; Variety (cybernetics); Strengths and weaknesses; Artificial intelligence; Machine translation; Machine learning; Data science; Benchmarking; Generative grammar; Ground truth; Natural language processing","score_opus":0.09181622296375748,"score_gpt":0.3453522069405535,"score_spread":0.25353598397679605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570371","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33298925,0.023709048,0.271501,0.0068085594,0.0035388756,0.006482657,0.14627305,0.1793512,0.029346371],"genre_scores_gemma":[0.30816302,0.0034056432,0.26855254,0.0026182767,0.0005265574,0.00483622,0.39824253,0.006295273,0.00735991],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9826034,0.009935427,0.0012842013,0.003198586,0.002463438,0.0005149803],"domain_scores_gemma":[0.95838505,0.02548428,0.0012202859,0.008475893,0.0051630624,0.0012714416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014533352,0.003927973,0.00207714,0.0049480633,0.0022200348,0.0035186273,0.005858508,0.002990917,0.005283283],"category_scores_gemma":[0.06078595,0.00093985116,0.0025812269,0.0048830095,0.0015237696,0.008354672,0.005137844,0.0051528425,0.005728723],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021455067,0.0022259713,0.026511684,0.011380426,0.001848899,0.00079542975,0.0022477624,0.08504689,0.011689022,0.0076661184,0.3557232,0.49271917],"study_design_scores_gemma":[0.0009795175,0.0028539272,0.028562512,0.0016599867,0.0008075238,0.0018524576,0.0036575166,0.71640456,0.027900843,0.0157899,0.19903208,0.00049926736],"about_ca_topic_score_codex":0.013811294,"about_ca_topic_score_gemma":0.025336716,"teacher_disagreement_score":0.014533352,"about_ca_system_score_codex":0.0033634477,"about_ca_system_score_gemma":0.003583676,"threshold_uncertainty_score":0.07686067},"labels":[],"label_agreement":null},{"id":"W4385570414","doi":"10.18653/v1/2023.acl-short.120","title":"Prefix Propagation: Parameter-Efficient Tuning for Long Sequences","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Prefix; Computer science; Sequence (biology); Focus (optics); Bridge (graph theory); Kernel (algebra); Language model; Simple (philosophy); Artificial intelligence; Algorithm; Theoretical computer science; Mathematics","score_opus":0.06876855535783585,"score_gpt":0.2941507825616755,"score_spread":0.22538222720383963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570414","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037873123,0.0007664047,0.9453422,0.00030572817,0.00017243701,0.000161431,0.00031705436,0.013207362,0.0018544341],"genre_scores_gemma":[0.55886525,0.00060572923,0.43066823,0.0005963764,0.0001565328,0.0006332331,0.0020561009,0.002268481,0.004150049],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9987949,0.00038890415,0.0000838625,0.0004412709,0.00016412293,0.00012695313],"domain_scores_gemma":[0.99688196,0.001715657,0.0001562195,0.0007339131,0.000375021,0.0001371917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026284482,0.0018782761,0.0010410973,0.0009297543,0.0006258025,0.0012877567,0.002725168,0.001822351,0.0060489667],"category_scores_gemma":[0.01569058,0.00084881973,0.000832885,0.0010926254,0.00093188894,0.0047770105,0.0024492505,0.004009985,0.003282448],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007106818,0.00046814707,0.00398826,0.00034526145,0.00024800914,0.00024754624,0.0004302158,0.37366155,0.022708677,0.010815988,0.012273452,0.5741022],"study_design_scores_gemma":[0.000040101975,0.00007206064,0.00027711136,0.000021559432,0.000023976923,0.000059867667,0.00003641259,0.9818463,0.005463395,0.010140447,0.001996161,0.000022636039],"about_ca_topic_score_codex":0.005142382,"about_ca_topic_score_gemma":0.00800711,"teacher_disagreement_score":0.0060489667,"about_ca_system_score_codex":0.0008943591,"about_ca_system_score_gemma":0.0021244658,"threshold_uncertainty_score":0.020235837},"labels":[],"label_agreement":null},{"id":"W4385570419","doi":"10.18653/v1/2023.americasnlp-1.10","title":"Towards the First Named Entity Recognition of Inuktitut for an Improved Machine Translation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Recall; Natural language processing; Artificial intelligence; Machine translation; Translation (biology); Indigenous; Named-entity recognition; Precision and recall; Word (group theory); Natural language; Speech recognition; Linguistics; Engineering","score_opus":0.0940507350730992,"score_gpt":0.28887577078290494,"score_spread":0.19482503570980575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570419","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11138281,0.0029518318,0.8445387,0.0014404971,0.00072913244,0.00034084014,0.005481646,0.014679266,0.01845522],"genre_scores_gemma":[0.47095796,0.0017384188,0.47675335,0.00041902327,0.00014496947,0.00023018835,0.027642658,0.0016768812,0.020436576],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99918216,0.0002243653,0.00005771813,0.00030470235,0.00012375701,0.000107192274],"domain_scores_gemma":[0.99876535,0.00019746396,0.000064177264,0.00027692513,0.0006404129,0.000055659835],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008568509,0.0011728927,0.0008120404,0.0017874441,0.0014814459,0.0021977066,0.0010376193,0.00082679826,0.004047985],"category_scores_gemma":[0.0022877608,0.00037279812,0.0011360866,0.0022064142,0.00046245192,0.0032379662,0.0015359371,0.0013164055,0.00537794],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037945216,0.000280541,0.012679366,0.0009396207,0.00030164665,0.0012737332,0.0022271054,0.04153939,0.062414456,0.034032628,0.059514448,0.7844177],"study_design_scores_gemma":[0.000039484225,0.00012904985,0.011496808,0.00018568529,0.0003693594,0.0011243061,0.0017492605,0.7590026,0.08164067,0.017943112,0.1261287,0.00019089837],"about_ca_topic_score_codex":0.10091051,"about_ca_topic_score_gemma":0.12672418,"teacher_disagreement_score":0.10091051,"about_ca_system_score_codex":0.0016558148,"about_ca_system_score_gemma":0.003896804,"threshold_uncertainty_score":0.20064628},"labels":[],"label_agreement":null},{"id":"W4385570658","doi":"10.18653/v1/2023.findings-acl.858","title":"Exploring the Effectiveness of Prompt Engineering for Legal Reasoning Tasks","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Task (project management); Computer science; Textual entailment; Natural language processing; Artificial intelligence; Cluster analysis; Logical consequence; Shot (pellet); Zero (linguistics); Best practice; Natural language; Linguistics; Engineering","score_opus":0.06120192517994883,"score_gpt":0.2559378797659381,"score_spread":0.1947359545859893,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570658","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6664236,0.00415181,0.258067,0.0020449134,0.00038466117,0.00071852526,0.0015499799,0.059700433,0.0069590765],"genre_scores_gemma":[0.83565897,0.0004414833,0.15891606,0.00036740594,0.00007977517,0.00016416234,0.0027239404,0.0003655958,0.0012825776],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9937264,0.0033827783,0.0004541809,0.0014518128,0.000754113,0.00023073514],"domain_scores_gemma":[0.9572973,0.03582625,0.0013297766,0.0029858332,0.0017311878,0.0008295465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010848443,0.0018921692,0.0010277658,0.002169505,0.0007425599,0.0017100794,0.0019390318,0.0018891521,0.0026842118],"category_scores_gemma":[0.060197942,0.00067029335,0.0007575174,0.000970465,0.00066983484,0.008042772,0.002211946,0.0031856194,0.0013052151],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037387288,0.0024480922,0.02311364,0.0014430581,0.00036719756,0.00034466013,0.0015691146,0.1160422,0.028676786,0.003925805,0.014357143,0.8039735],"study_design_scores_gemma":[0.00034958805,0.001514392,0.0057150545,0.000079563135,0.00015790248,0.00023149437,0.0004990008,0.9522281,0.025726136,0.008165303,0.005241294,0.00009218068],"about_ca_topic_score_codex":0.007514655,"about_ca_topic_score_gemma":0.009684598,"teacher_disagreement_score":0.010848443,"about_ca_system_score_codex":0.0016445684,"about_ca_system_score_gemma":0.0022948058,"threshold_uncertainty_score":0.05737275},"labels":[],"label_agreement":null},{"id":"W4385570721","doi":"10.18653/v1/2023.findings-acl.625","title":"EmbedTextNet: Dimension Reduction with Weighted Reconstruction and Correlation Losses for Efficient Text Embedding","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Embedding; Computer science; Language model; Benchmark (surveying); Similarity (geometry); Latency (audio); Reduction (mathematics); Dimension (graph theory); Correlation; Dimensionality reduction; Code (set theory); Source lines of code; Artificial intelligence; Machine learning; Pattern recognition (psychology); Data mining; Mathematics; Programming language","score_opus":0.016322133604309903,"score_gpt":0.2510149713938844,"score_spread":0.2346928377895745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570721","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02217224,0.0003760885,0.97003114,0.00025120168,0.00011193709,0.00013293125,0.00048058064,0.0048583443,0.0015855056],"genre_scores_gemma":[0.341053,0.0006693483,0.64203024,0.0003161443,0.00016348097,0.0008028174,0.004027038,0.0011211378,0.009816748],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995449,0.00014184923,0.000032040214,0.000105548395,0.00012442097,0.00005115668],"domain_scores_gemma":[0.9988011,0.0005388494,0.000094061426,0.00033549368,0.00017181583,0.000058655023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010826106,0.0014979488,0.0007674993,0.00090942206,0.0003865823,0.0009072061,0.001601523,0.0011176468,0.003928488],"category_scores_gemma":[0.0054499856,0.00046839035,0.00079729955,0.0008901762,0.0006457068,0.0036604756,0.0021369755,0.0018170953,0.002643421],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041486684,0.00031314566,0.0014595516,0.00040851944,0.00016407449,0.00038452502,0.00029417578,0.4334654,0.023203975,0.029509485,0.020447014,0.48993528],"study_design_scores_gemma":[0.000018406396,0.000051520114,0.00008303502,0.000010729325,0.000012024091,0.000068309906,0.000025026373,0.9833243,0.0046030083,0.009373848,0.00241787,0.000011857482],"about_ca_topic_score_codex":0.0024647017,"about_ca_topic_score_gemma":0.005492519,"teacher_disagreement_score":0.003928488,"about_ca_system_score_codex":0.00068642944,"about_ca_system_score_gemma":0.000872405,"threshold_uncertainty_score":0.013142109},"labels":[],"label_agreement":null},{"id":"W4385570792","doi":"10.18653/v1/2023.findings-acl.558","title":"MVP: Multi-task Supervised Pre-training for Natural Language Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Natural Science Foundation of China","keywords":"Computer science; Natural language generation; Task (project management); Artificial intelligence; Generality; Natural language processing; Language model; Natural language; Machine learning; Scale (ratio); Natural language understanding; Supervised learning; Artificial neural network","score_opus":0.08159341295343946,"score_gpt":0.31775254635337474,"score_spread":0.23615913339993527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570792","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039595317,0.0018630247,0.8916691,0.0007007796,0.000489669,0.0006603799,0.002771723,0.05666924,0.0055806935],"genre_scores_gemma":[0.3749739,0.0005667025,0.5920927,0.0014529544,0.00023738867,0.0020381678,0.015877424,0.0028356593,0.009925126],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99876887,0.00047759386,0.000062958214,0.00042478222,0.0001636462,0.00010209571],"domain_scores_gemma":[0.9972414,0.0017104007,0.000121878,0.00040838952,0.0004067832,0.00011120934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002041246,0.0022203878,0.00096965407,0.0009085096,0.0006047956,0.00074937183,0.002847737,0.0016949808,0.006611063],"category_scores_gemma":[0.005679128,0.00080616574,0.0013835989,0.0008437097,0.00067279575,0.0019477255,0.0018460714,0.003989958,0.0037401752],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007116435,0.0006355279,0.002519499,0.0006499621,0.00021698189,0.00031890272,0.00040806585,0.18791349,0.018412521,0.0037890195,0.057395235,0.72702926],"study_design_scores_gemma":[0.00010089378,0.00019766086,0.0005566767,0.000035542038,0.000032478234,0.00008602056,0.000053010794,0.97533786,0.011362108,0.0051473896,0.0070624594,0.00002795066],"about_ca_topic_score_codex":0.0057656644,"about_ca_topic_score_gemma":0.01341588,"teacher_disagreement_score":0.006611063,"about_ca_system_score_codex":0.0009152465,"about_ca_system_score_gemma":0.0017286523,"threshold_uncertainty_score":0.022116184},"labels":[],"label_agreement":null},{"id":"W4385570795","doi":"10.18653/v1/2023.findings-acl.526","title":"Open-WikiTable : Dataset for Open Domain Question Answering with Complex Reasoning over Table","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Samsung; Korea Health Industry Development Institute; National Research Foundation of Korea; Ministry of Science and ICT, South Korea; National Research Foundation","keywords":"Open domain; Computer science; Question answering; Parsing; Table (database); Domain (mathematical analysis); Open research; Task (project management); SQL; Range (aeronautics); Information retrieval; Sorting; Artificial intelligence; Data mining; World Wide Web; Programming language","score_opus":0.06894793121732369,"score_gpt":0.34013305711222724,"score_spread":0.27118512589490357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570795","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048097325,0.0005473452,0.0041438094,0.000556428,0.00012455808,0.00016934569,0.9790758,0.0062337043,0.004339121],"genre_scores_gemma":[0.0049564047,0.000116730276,0.005858715,0.000121950456,0.000019478364,0.00018212773,0.9875907,0.00023560248,0.00091822696],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99678254,0.0007500748,0.0005502997,0.00089151395,0.0007583609,0.00026721874],"domain_scores_gemma":[0.9932715,0.002690417,0.000619911,0.0016056582,0.0011819905,0.0006305791],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018121877,0.002544454,0.0011131841,0.004754243,0.0013004185,0.0025614458,0.003311236,0.0025972975,0.015762117],"category_scores_gemma":[0.014529734,0.00064222416,0.001941681,0.006120942,0.00077290786,0.00447203,0.0039258394,0.0024252718,0.01839682],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022339768,0.00019076279,0.005079104,0.0021993727,0.000089154135,0.00019312819,0.00043570346,0.0022130404,0.001132983,0.0055502597,0.96375376,0.018939357],"study_design_scores_gemma":[0.00030806428,0.00008369308,0.010212536,0.00056388223,0.00006465012,0.00046929996,0.0012454657,0.012650102,0.0031231917,0.0141329495,0.9570311,0.000115124974],"about_ca_topic_score_codex":0.028131213,"about_ca_topic_score_gemma":0.055945095,"teacher_disagreement_score":0.028131213,"about_ca_system_score_codex":0.002294914,"about_ca_system_score_gemma":0.003375648,"threshold_uncertainty_score":0.055934966},"labels":[],"label_agreement":null},{"id":"W4385570798","doi":"10.18653/v1/2023.acl-long.136","title":"Augmentation-Adapted Retriever Improves Generalization of Language Models as Generic Plug-In","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Tsinghua University","keywords":"Generalization; Computer science; Scheme (mathematics); Source code; Labrador Retriever; Open source; Plug-in; Set (abstract data type); Plug and play; Code (set theory); Artificial intelligence; Operating system; Mathematics; Programming language","score_opus":0.031147238484415005,"score_gpt":0.2720054915364587,"score_spread":0.2408582530520437,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570798","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.101410106,0.003563711,0.7171318,0.00076971087,0.0005357227,0.00035247696,0.0035960982,0.16083352,0.01180687],"genre_scores_gemma":[0.5265075,0.0015049408,0.41285795,0.0018776306,0.00034788184,0.00071903045,0.02033209,0.00824588,0.027607],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986368,0.00041144626,0.00008062786,0.0005132306,0.0002244127,0.00013356176],"domain_scores_gemma":[0.99841106,0.00054996717,0.00006311736,0.0006682318,0.00024170012,0.00006596167],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018000253,0.0026276757,0.0014661965,0.0009594684,0.00046659182,0.0013278073,0.002277973,0.0013958493,0.008124816],"category_scores_gemma":[0.0058510923,0.00069226156,0.0019073025,0.00077233097,0.00065427774,0.003847069,0.002583617,0.0024054893,0.012023636],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091981224,0.00058495265,0.0022206684,0.0007422203,0.00034362153,0.0007015707,0.000516113,0.10464545,0.082328595,0.003712394,0.055430807,0.7478539],"study_design_scores_gemma":[0.00012974151,0.0003664499,0.0011357173,0.000044386703,0.00014147439,0.00045925475,0.00018913481,0.92684877,0.04613468,0.005847133,0.018602693,0.000100602796],"about_ca_topic_score_codex":0.004947495,"about_ca_topic_score_gemma":0.008355929,"teacher_disagreement_score":0.008124816,"about_ca_system_score_codex":0.00058522745,"about_ca_system_score_gemma":0.0008410865,"threshold_uncertainty_score":0.027180254},"labels":[],"label_agreement":null},{"id":"W4385570832","doi":"10.18653/v1/2023.findings-acl.896","title":"DEnsity: Open-domain Dialogue Evaluation Metric using Density Estimation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Metric (unit); Computer science; Classifier (UML); Artificial intelligence; Feature vector; Feature (linguistics); Machine learning; Density estimation; Domain (mathematical analysis); Open domain; Pattern recognition (psychology); Mathematics; Statistics","score_opus":0.13448994961271787,"score_gpt":0.3537995330139125,"score_spread":0.2193095834011946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570832","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12189823,0.0032996018,0.8581152,0.00058741483,0.00026198936,0.0007076108,0.002256919,0.005415079,0.0074579325],"genre_scores_gemma":[0.86677617,0.0004418031,0.12647963,0.00018266712,0.00018127015,0.0007644724,0.0027722167,0.00038319267,0.0020184566],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98875695,0.0056025623,0.0008723614,0.0014226702,0.0029573715,0.00038812993],"domain_scores_gemma":[0.9678083,0.020881532,0.0024296457,0.0022602642,0.0056574796,0.00096278137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007269303,0.0018405118,0.0013693939,0.0036038957,0.0006382561,0.0018480469,0.0013053766,0.0017886295,0.0019768472],"category_scores_gemma":[0.048246384,0.0002960192,0.0007178597,0.0015687877,0.0010748816,0.0040078564,0.0026043581,0.0017289954,0.0012187284],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026857918,0.0010066078,0.05577647,0.0017385898,0.0006817944,0.00044573733,0.0022726632,0.1569954,0.020902278,0.020387907,0.024874438,0.7122324],"study_design_scores_gemma":[0.00009656521,0.0010416529,0.023920136,0.0001715556,0.0001461205,0.0007510281,0.0006537737,0.9233988,0.014815987,0.02665351,0.008101658,0.00024929387],"about_ca_topic_score_codex":0.0026517212,"about_ca_topic_score_gemma":0.0019010043,"teacher_disagreement_score":0.007269303,"about_ca_system_score_codex":0.0013378934,"about_ca_system_score_gemma":0.0009340593,"threshold_uncertainty_score":0.03844422},"labels":[],"label_agreement":null},{"id":"W4385570834","doi":"10.18653/v1/2023.acl-long.102","title":"Few-shot Adaptation Works with UnpredicTable Data","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute for Catastrophic Loss Reduction","keywords":"Computer science; Adaptation (eye); Domain adaptation; Machine learning; Artificial intelligence; Documentation; Classifier (UML)","score_opus":0.14374101805130748,"score_gpt":0.28447914128191604,"score_spread":0.14073812323060855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570834","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25999698,0.0058542597,0.6935286,0.0027253148,0.0012762359,0.0004894144,0.002881599,0.020684706,0.0125627965],"genre_scores_gemma":[0.8554865,0.0007974637,0.12185464,0.002472319,0.0004230555,0.00040353605,0.008259923,0.0013417021,0.008960885],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9967051,0.0008707051,0.00016082435,0.0016584863,0.00039065245,0.00021410664],"domain_scores_gemma":[0.99046004,0.005525983,0.00027348593,0.0027701112,0.0006438985,0.00032656247],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046242652,0.0020408614,0.001940855,0.0011952155,0.0011197509,0.002077649,0.0030074252,0.0020021535,0.0031973089],"category_scores_gemma":[0.02400542,0.00093416165,0.0011822158,0.00096364104,0.0015045961,0.005720746,0.0028779102,0.004548963,0.0030976848],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012310942,0.001369851,0.021667145,0.0011309206,0.00087736285,0.0005977274,0.0008994186,0.37534237,0.0319276,0.00650879,0.039255418,0.5191924],"study_design_scores_gemma":[0.000068013695,0.00027877314,0.003834515,0.00008111631,0.000102805694,0.00026020623,0.00025875092,0.95612407,0.010237413,0.021835787,0.006830972,0.00008754845],"about_ca_topic_score_codex":0.006016859,"about_ca_topic_score_gemma":0.009970161,"teacher_disagreement_score":0.006016859,"about_ca_system_score_codex":0.0012091635,"about_ca_system_score_gemma":0.001190783,"threshold_uncertainty_score":0.024455726},"labels":[],"label_agreement":null},{"id":"W4385570852","doi":"10.18653/v1/2023.acl-long.201","title":"MatSci-NLP: Evaluating Scientific Language Models on Materials Science Language Tasks Using Text-to-Schema Modeling","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Artificial intelligence; Natural language processing; Schema (genetic algorithms); Benchmark (surveying); Language model; Question answering; Information retrieval","score_opus":0.12813830277042124,"score_gpt":0.3768265016401982,"score_spread":0.24868819886977697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570852","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4202465,0.012503873,0.21981558,0.0071320753,0.003663641,0.0035114477,0.13165063,0.15481229,0.04666398],"genre_scores_gemma":[0.44593173,0.0016710195,0.19033726,0.0027597393,0.00041863744,0.0035293463,0.33997726,0.0036756592,0.011699254],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99583745,0.0016288124,0.00039084617,0.0011621562,0.0006996458,0.0002811194],"domain_scores_gemma":[0.98582804,0.009654353,0.0005315399,0.001906559,0.0014836568,0.0005957991],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006383211,0.0044014314,0.0013438035,0.0027658131,0.0012214128,0.002841688,0.0046965596,0.003997944,0.01318918],"category_scores_gemma":[0.02502784,0.0009528665,0.0024480845,0.0027607647,0.0015477061,0.005641138,0.003247491,0.0055328477,0.008361138],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025510183,0.0030520977,0.010322268,0.0050606932,0.0011337394,0.00082653743,0.0006415801,0.36027557,0.015021559,0.006393083,0.23527291,0.359449],"study_design_scores_gemma":[0.00073229335,0.0010770299,0.004741131,0.00021119171,0.00018227962,0.00028714622,0.00051650597,0.9296927,0.020842336,0.011404216,0.030172648,0.00014061546],"about_ca_topic_score_codex":0.019088939,"about_ca_topic_score_gemma":0.021640148,"teacher_disagreement_score":0.019088939,"about_ca_system_score_codex":0.00290646,"about_ca_system_score_gemma":0.0041826908,"threshold_uncertainty_score":0.04412216},"labels":[],"label_agreement":null},{"id":"W4385570971","doi":"10.18653/v1/2023.findings-acl.169","title":"From chocolate bunny to chocolate crocodile: Do Language Models Understand Noun Compounds?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of British Columbia","funders":"Vector Institute; Universities Space Research Association; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Noun; Paraphrase; Task (project management); Interpretation (philosophy); Natural language processing; Conceptualization; Computer science; Artificial intelligence; Linguistics; Philosophy; Programming language","score_opus":0.038597233571667604,"score_gpt":0.27824413171395684,"score_spread":0.23964689814228923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385570971","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6177689,0.0019408282,0.32083333,0.00666005,0.0005526467,0.00036401342,0.011395907,0.009364762,0.031119667],"genre_scores_gemma":[0.87805647,0.00045439482,0.1016521,0.0010260508,0.00009626599,0.00013575084,0.014565345,0.00070251833,0.0033111186],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99915004,0.00033503238,0.000044477187,0.00030802703,0.00009486827,0.00006762262],"domain_scores_gemma":[0.99475694,0.0037025597,0.00025543215,0.00071679405,0.00041346793,0.00015481534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016835212,0.0011804755,0.0005620533,0.00076049555,0.0005034716,0.0028489341,0.0015033782,0.0016594424,0.006320489],"category_scores_gemma":[0.011181964,0.00041616708,0.0013686012,0.0008355507,0.0006068673,0.0071115103,0.0015465512,0.0030299774,0.0030976434],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00178935,0.0006084567,0.05732978,0.0021768487,0.0007963898,0.0019484964,0.0064842254,0.10264655,0.05900413,0.027717637,0.07546082,0.6640373],"study_design_scores_gemma":[0.00019129667,0.00017179936,0.014799082,0.00023473344,0.00015333697,0.00083818805,0.003623983,0.85698533,0.01578414,0.08047797,0.026631795,0.000108288456],"about_ca_topic_score_codex":0.00983605,"about_ca_topic_score_gemma":0.012747018,"teacher_disagreement_score":0.00983605,"about_ca_system_score_codex":0.0010475662,"about_ca_system_score_gemma":0.001206361,"threshold_uncertainty_score":0.021144092},"labels":[],"label_agreement":null},{"id":"W4385571010","doi":"10.18653/v1/2023.acl-demo.57","title":"GAIA Search: Hugging Face and Pyserini Interoperability for NLP Training Data Exploration","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Interoperability; STELLA (programming language); Natural language processing; Artificial intelligence; Training set; Information retrieval; Face (sociological concept); World Wide Web; Linguistics; Philosophy","score_opus":0.5397407610899357,"score_gpt":0.3876501314400893,"score_spread":0.15209062964984643,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571010","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01729388,0.0020001542,0.48453763,0.0023106406,0.0009281798,0.0003782498,0.009772996,0.4639393,0.018838977],"genre_scores_gemma":[0.19061461,0.001372535,0.6955846,0.002193152,0.0004125114,0.0011482032,0.049823523,0.029776538,0.029074265],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982974,0.00061206205,0.000098446915,0.00047052183,0.0003662452,0.00015528368],"domain_scores_gemma":[0.99775094,0.0011824862,0.000044495624,0.00068726257,0.00018610455,0.00014869164],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003119764,0.0015020088,0.0014710426,0.0019930464,0.001506759,0.0029093497,0.0032504129,0.0025399283,0.029268961],"category_scores_gemma":[0.008422912,0.0008645592,0.0013590067,0.0021518285,0.0008940072,0.007555256,0.0077058016,0.0023815234,0.021667141],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002076013,0.00032979692,0.0015784086,0.00040641785,0.00020899222,0.00041113683,0.0010602291,0.0070014577,0.009173021,0.016257731,0.40903217,0.5524646],"study_design_scores_gemma":[0.0007047473,0.0003973615,0.0019078729,0.0001868898,0.00017610197,0.0007254737,0.0011991254,0.5786635,0.03138105,0.07437218,0.31007463,0.00021100962],"about_ca_topic_score_codex":0.0060084295,"about_ca_topic_score_gemma":0.0116002215,"teacher_disagreement_score":0.029268961,"about_ca_system_score_codex":0.0006702298,"about_ca_system_score_gemma":0.0013675289,"threshold_uncertainty_score":0.0979144},"labels":[],"label_agreement":null},{"id":"W4385571031","doi":"10.18653/v1/2023.findings-acl.868","title":"NusaCrowd: Open Source Initiative for Indonesian NLP Resources","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Indonesian; Art; Artificial intelligence; Art history; Theology; Philosophy; Computer science; Linguistics","score_opus":0.08243536302076786,"score_gpt":0.31385751820298324,"score_spread":0.23142215518221537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571031","genre_codex":"software","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065607927,0.0029112718,0.20593484,0.0054057445,0.0019340235,0.0010210812,0.30426905,0.3949553,0.077007905],"genre_scores_gemma":[0.031450897,0.0018763483,0.29443264,0.0014932115,0.0002797028,0.0024408435,0.544182,0.09318957,0.030654822],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9951799,0.0012845397,0.00076906587,0.0010795905,0.001427744,0.0002591746],"domain_scores_gemma":[0.98362964,0.007251895,0.0007657227,0.004716051,0.0018432722,0.001793401],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0067423773,0.0021699788,0.0019333472,0.007826923,0.0027818524,0.0064897668,0.004298312,0.0020923493,0.045058846],"category_scores_gemma":[0.02326572,0.002027586,0.0025968428,0.007533232,0.0019144372,0.013994932,0.013346913,0.0046483586,0.04841894],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056798966,0.0001842422,0.001960602,0.0029700806,0.00021603047,0.00072956434,0.0019699496,0.0018152865,0.004217893,0.017293187,0.8133937,0.1546815],"study_design_scores_gemma":[0.00012172778,0.000028307455,0.002554084,0.00048504246,0.000056932604,0.0003900977,0.0007850651,0.011741336,0.00452856,0.019264037,0.95989394,0.00015087386],"about_ca_topic_score_codex":0.011277722,"about_ca_topic_score_gemma":0.011563615,"teacher_disagreement_score":0.9957017,"about_ca_system_score_codex":0.002039939,"about_ca_system_score_gemma":0.0064166556,"threshold_uncertainty_score":0.15073687},"labels":[],"label_agreement":null},{"id":"W4385571080","doi":"10.18653/v1/2023.acl-long.508","title":"Grounded Multimodal Named Entity Recognition on Social Media","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Fields Institute for Research in Mathematical Sciences","funders":"Government of Jiangsu Province; Natural Science Foundation of Jiangsu Province; National Natural Science Foundation of China","keywords":"Computer science; Baseline (sea); Task (project management); Construct (python library); Entity linking; Artificial intelligence; Bounding overwatch; Information retrieval; Named-entity recognition; Social media; Natural language processing; Index (typography); Graph; World Wide Web; Knowledge base; Theoretical computer science","score_opus":0.10155389050067026,"score_gpt":0.28991084789368937,"score_spread":0.1883569573930191,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571080","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13764958,0.0072085955,0.6872911,0.002513331,0.00078973925,0.0012316386,0.0891001,0.05415851,0.020057438],"genre_scores_gemma":[0.391736,0.0019196226,0.40635356,0.00093514344,0.00046558568,0.00093275204,0.1844338,0.0009328031,0.012290741],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99736834,0.0006917293,0.00023667993,0.0010667976,0.0004562086,0.00018023614],"domain_scores_gemma":[0.9958858,0.0013130667,0.00045412293,0.0015756659,0.0006125491,0.00015874523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017998128,0.0020958271,0.0012112008,0.0064172973,0.0012378473,0.0018187534,0.0020014618,0.0020304113,0.0047506527],"category_scores_gemma":[0.0068807695,0.00041760507,0.0016334793,0.0046660304,0.0006713709,0.009388318,0.0036426678,0.0014284968,0.004658956],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008222702,0.00057109335,0.0120110605,0.0015514988,0.00042849962,0.00181432,0.0010166509,0.02736272,0.030388772,0.022531739,0.13116246,0.770339],"study_design_scores_gemma":[0.00006463966,0.0002686849,0.016178243,0.00027261823,0.0003179085,0.0017249354,0.0022689207,0.6821108,0.069268145,0.05765288,0.16963379,0.00023845767],"about_ca_topic_score_codex":0.006675407,"about_ca_topic_score_gemma":0.013029802,"teacher_disagreement_score":0.006675407,"about_ca_system_score_codex":0.0010971957,"about_ca_system_score_gemma":0.00079675805,"threshold_uncertainty_score":0.015892446},"labels":[],"label_agreement":null},{"id":"W4385571088","doi":"10.18653/v1/2023.acl-short.89","title":"Decomposed scoring of CCG dependencies","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Natural language processing; Predicate (mathematical logic); Parsing; Artificial intelligence; Ranking (information retrieval); Judgement; Grammar; Argument (complex analysis); Task (project management); Linguistics; Programming language","score_opus":0.04587723139958526,"score_gpt":0.270129075844673,"score_spread":0.22425184444508772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571088","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20843714,0.00070607057,0.7574592,0.0008467471,0.000532953,0.0008761397,0.0018565325,0.010411032,0.01887414],"genre_scores_gemma":[0.686196,0.00015262532,0.30188522,0.00032963956,0.00010036693,0.00045924453,0.002757365,0.002497379,0.005622059],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96753705,0.017180229,0.0019820756,0.004930679,0.0070569,0.0013129719],"domain_scores_gemma":[0.9072157,0.045266226,0.003959793,0.014400628,0.027159426,0.0019982434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026284361,0.0022651465,0.0018944867,0.0055809785,0.0017688695,0.0051085064,0.0026075842,0.0021982747,0.0069708857],"category_scores_gemma":[0.099545784,0.00075611495,0.0008594405,0.0044204616,0.0018894891,0.0034622422,0.004218185,0.0029149267,0.0030765557],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012967533,0.00032729795,0.04925729,0.0007514508,0.00042928936,0.0005957592,0.00538027,0.048300236,0.047950283,0.033968024,0.03387659,0.77786666],"study_design_scores_gemma":[0.00019001855,0.0006451592,0.047236685,0.00029019266,0.00024513688,0.0008910034,0.002334752,0.78262043,0.05432665,0.07691321,0.03382458,0.00048213376],"about_ca_topic_score_codex":0.005489251,"about_ca_topic_score_gemma":0.012193047,"teacher_disagreement_score":0.026284361,"about_ca_system_score_codex":0.002047362,"about_ca_system_score_gemma":0.0036221086,"threshold_uncertainty_score":0.13900667},"labels":[],"label_agreement":null},{"id":"W4385571112","doi":"10.18653/v1/2023.findings-acl.256","title":"Search-Oriented Conversational Query Editing","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Ministry of Education, India; Renmin University of China; National Natural Science Foundation of China","keywords":"Computer science; Security token; Contextualization; Rewriting; Ranking (information retrieval); Search engine; Information retrieval; Artificial intelligence; Natural language processing; Programming language","score_opus":0.039871876088724216,"score_gpt":0.27339322807950317,"score_spread":0.23352135199077895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571112","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043781623,0.0011836747,0.93607146,0.00056892773,0.00013643339,0.000472489,0.00074722053,0.008557384,0.0084808925],"genre_scores_gemma":[0.7318789,0.00053830096,0.25161257,0.0005475356,0.00016251103,0.00040511013,0.0015203796,0.0007132179,0.012621509],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99782205,0.000879017,0.00014141417,0.00057797134,0.00044754086,0.00013206164],"domain_scores_gemma":[0.9972262,0.0014134422,0.00015971456,0.0005024159,0.00057850656,0.000119720426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001993996,0.0010563975,0.0010776725,0.00064514857,0.0004823959,0.001177565,0.002444289,0.000948058,0.0044115707],"category_scores_gemma":[0.007703574,0.00034700357,0.0009746017,0.00060300407,0.0007695459,0.0023327176,0.0018186589,0.0013325264,0.0018190718],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015052017,0.0005291813,0.0031665908,0.0014290403,0.00025883227,0.000932885,0.0031332944,0.18930799,0.15647446,0.050919633,0.028038858,0.56430405],"study_design_scores_gemma":[0.000060054226,0.00018972592,0.00047538054,0.000019500885,0.00008577918,0.00025635664,0.00017977644,0.9479845,0.024800375,0.011396944,0.014491503,0.000060011353],"about_ca_topic_score_codex":0.0053593926,"about_ca_topic_score_gemma":0.0052460255,"teacher_disagreement_score":0.0053593926,"about_ca_system_score_codex":0.00061633537,"about_ca_system_score_gemma":0.0011863088,"threshold_uncertainty_score":0.01475817},"labels":[],"label_agreement":null},{"id":"W4385571149","doi":"10.18653/v1/2023.findings-acl.211","title":"Exploiting Hierarchically Structured Categories in Fine-grained Chinese Named Entity Recognition","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Relevance (law); Artificial intelligence; Function (biology); Natural language processing; Information retrieval; Pattern recognition (psychology)","score_opus":0.026900394584444962,"score_gpt":0.26114673704919533,"score_spread":0.23424634246475037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571149","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30074894,0.008239631,0.59784514,0.0016862045,0.000593867,0.0014324313,0.03710035,0.031012962,0.021340488],"genre_scores_gemma":[0.6293663,0.0011866823,0.29455346,0.00065432524,0.00017533769,0.00041759165,0.06493062,0.00035722766,0.008358502],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998078,0.0003841553,0.00018108981,0.00081840344,0.0003244621,0.00021392119],"domain_scores_gemma":[0.9964888,0.0010515453,0.00033342728,0.0012422956,0.0007094857,0.0001743543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023969805,0.0012772408,0.0011476956,0.0061525283,0.0011746164,0.0011980166,0.0018322982,0.0012898189,0.0028344805],"category_scores_gemma":[0.0053503215,0.0003245435,0.0011469426,0.0046887132,0.00065439,0.0049578864,0.0020901323,0.0014412572,0.0021406107],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065218675,0.0004584661,0.03381131,0.0010943756,0.00025484335,0.0012661396,0.0011260957,0.033133827,0.032232735,0.018113531,0.058868885,0.8189877],"study_design_scores_gemma":[0.00012964515,0.0004070184,0.035868358,0.0002194816,0.0002906073,0.0013714979,0.0012752186,0.77553964,0.043475837,0.047143836,0.09401776,0.00026107373],"about_ca_topic_score_codex":0.022986729,"about_ca_topic_score_gemma":0.047352888,"teacher_disagreement_score":0.022986729,"about_ca_system_score_codex":0.0011721993,"about_ca_system_score_gemma":0.0018600523,"threshold_uncertainty_score":0.045705855},"labels":[],"label_agreement":null},{"id":"W4385571251","doi":"10.18653/v1/2023.acl-long.317","title":"Simplicity Bias in Transformers and their Ability to Learn Sparse Boolean Functions","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"University of Oxford","keywords":"Transformer; Overfitting; Computer science; Inductive bias; Generalization; Boolean function; Machine learning; Artificial intelligence; Mathematics; Algorithm; Multi-task learning; Electrical engineering; Engineering; Artificial neural network","score_opus":0.0824883815686029,"score_gpt":0.27474028526736105,"score_spread":0.19225190369875816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571251","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27414143,0.000640506,0.71425587,0.0013952018,0.00006746983,0.00008932395,0.0004154233,0.002722658,0.0062721516],"genre_scores_gemma":[0.96172327,0.00032543187,0.035841234,0.00023370462,0.000033053715,0.000049492704,0.0003240607,0.00018753746,0.0012821581],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987458,0.0005419896,0.00008768107,0.00029184882,0.00022829334,0.000104468396],"domain_scores_gemma":[0.9822081,0.014573233,0.00084841595,0.0015630026,0.0005675894,0.00023958541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005135378,0.00080442964,0.000730376,0.0007586107,0.0003007336,0.001632757,0.0010612801,0.00097533874,0.0025077772],"category_scores_gemma":[0.031059856,0.00053230056,0.0008689343,0.0005534893,0.0015064642,0.0060837846,0.0016236609,0.002367491,0.00075354095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086567237,0.0001663165,0.024658559,0.00087196863,0.0003440398,0.00040540178,0.0015380428,0.46729934,0.041095696,0.13576046,0.0048902305,0.32210433],"study_design_scores_gemma":[0.0000281806,0.00019747563,0.0018087706,0.00004189885,0.000064469685,0.00028402556,0.000093116985,0.8784819,0.0095549375,0.108271964,0.0011422852,0.00003098636],"about_ca_topic_score_codex":0.0017958038,"about_ca_topic_score_gemma":0.002795618,"teacher_disagreement_score":0.005135378,"about_ca_system_score_codex":0.000876147,"about_ca_system_score_gemma":0.0006706914,"threshold_uncertainty_score":0.027158797},"labels":[],"label_agreement":null},{"id":"W4385571325","doi":"10.18653/v1/2023.acl-long.230","title":"A fine-grained comparison of pragmatic language understanding in humans and language models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Literal (mathematical logic); Computer science; Pragmatics; Heuristic; Set (abstract data type); Language model; Natural language processing; Artificial intelligence; Human language; Linguistics","score_opus":0.07989764519677654,"score_gpt":0.3256151878462655,"score_spread":0.24571754264948892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571325","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93847275,0.00013704687,0.056558926,0.00029637548,0.000019621053,0.00008832115,0.00024108232,0.00043470477,0.0037512456],"genre_scores_gemma":[0.98546314,0.000046575817,0.013807766,0.00006119257,0.000004723289,0.000045292207,0.00026511375,0.000052426763,0.00025386343],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9939778,0.003379643,0.00026571492,0.001316868,0.00085400895,0.00020609978],"domain_scores_gemma":[0.9701812,0.02220559,0.001678066,0.004268933,0.0011277037,0.00053846155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061006923,0.0006105062,0.0006064854,0.0008521196,0.00049327995,0.003236783,0.00077539653,0.0013057298,0.0015875847],"category_scores_gemma":[0.04243857,0.0005238039,0.0004884874,0.00034855359,0.002050514,0.0046210815,0.0025654386,0.0012588617,0.000394052],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024918283,0.0015261263,0.23522444,0.0018233187,0.00086285983,0.0009919369,0.083781555,0.10002012,0.282074,0.04305609,0.0051366813,0.243011],"study_design_scores_gemma":[0.00019022953,0.0025038682,0.20791201,0.00023504779,0.00021574361,0.0013545118,0.022731068,0.59353274,0.06742853,0.090888366,0.012476397,0.0005315287],"about_ca_topic_score_codex":0.0017618991,"about_ca_topic_score_gemma":0.0021046614,"teacher_disagreement_score":0.0061006923,"about_ca_system_score_codex":0.0006612149,"about_ca_system_score_gemma":0.00074275275,"threshold_uncertainty_score":0.032263994},"labels":[],"label_agreement":null},{"id":"W4385571397","doi":"10.18653/v1/2023.clinicalnlp-1.36","title":"WangLab at MEDIQA-Chat 2023: Clinical Note Generation from Doctor-Patient Conversations using Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Health Sciences Centre; University Health Network; Sunnybrook Health Science Centre; Vector Institute; University of Toronto","funders":"Alliance de recherche numérique du Canada","keywords":"Computer science; Task (project management); Scrutiny; Context (archaeology); Language model; Natural language processing; Path (computing); Artificial intelligence; Human–computer interaction; Multimedia; World Wide Web; Programming language","score_opus":0.11737923125211153,"score_gpt":0.3526829460646227,"score_spread":0.23530371481251117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571397","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16424133,0.0027105946,0.50898814,0.011174634,0.010654699,0.0058929706,0.07517458,0.20043583,0.020727238],"genre_scores_gemma":[0.34840336,0.00057988346,0.45765093,0.0026358992,0.0022853832,0.0038979885,0.13973074,0.015743846,0.029071886],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98680866,0.0074138115,0.0007106037,0.0021362999,0.0021643357,0.00076626835],"domain_scores_gemma":[0.9550875,0.022575477,0.0009911939,0.0072077136,0.008852476,0.005285581],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014386413,0.0027528573,0.0017376712,0.0016343427,0.0020917153,0.0030105107,0.0038708323,0.003816245,0.028130211],"category_scores_gemma":[0.046923324,0.0010548026,0.0017107454,0.00087164884,0.001034281,0.0024686968,0.006157951,0.004354811,0.022261357],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0065567004,0.0016992259,0.0057551544,0.0025233906,0.0005457092,0.0023563784,0.004303536,0.018069144,0.07032323,0.0026588175,0.5453212,0.3398874],"study_design_scores_gemma":[0.0049166526,0.005508635,0.020523062,0.00059604907,0.0003977232,0.004499616,0.0046543074,0.41886756,0.12999785,0.017079128,0.39161968,0.0013396627],"about_ca_topic_score_codex":0.0062857373,"about_ca_topic_score_gemma":0.009333034,"teacher_disagreement_score":0.028130211,"about_ca_system_score_codex":0.0014746969,"about_ca_system_score_gemma":0.003814907,"threshold_uncertainty_score":0.094104886},"labels":[],"label_agreement":null},{"id":"W4385571398","doi":"10.18653/v1/2023.acl-long.95","title":"Fine-tuning Happens in Tiny Subspaces: Exploring Intrinsic Task-specific Subspaces of Pre-trained Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Canadian Institute for Advanced Research","funders":"Fundamental Research Funds for the Central Universities; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China; Canadian Institute for Advanced Research","keywords":"Subspace topology; Computer science; Linear subspace; Parameterized complexity; Task (project management); Redundancy (engineering); Outlier; Artificial intelligence; Perspective (graphical); Process (computing); Machine learning; Algorithm; Mathematics","score_opus":0.08221896639818659,"score_gpt":0.2670756845490518,"score_spread":0.18485671815086518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571398","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33728927,0.00082423264,0.6576708,0.00048940605,0.0000653614,0.00006530002,0.00020123798,0.0016278228,0.0017665721],"genre_scores_gemma":[0.93227595,0.00024278768,0.06539173,0.00018494425,0.000037750633,0.000072986564,0.0005353579,0.00028279755,0.0009756186],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994837,0.0001703077,0.000023858947,0.00018921173,0.000060595783,0.00007230557],"domain_scores_gemma":[0.9980136,0.0011428305,0.00015587076,0.00047382762,0.000121708734,0.00009220226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001293405,0.0010557143,0.0008935567,0.00035778934,0.00038289058,0.0010035913,0.0008093466,0.0007596905,0.0008811984],"category_scores_gemma":[0.0090972325,0.00053377706,0.00074626965,0.00038430217,0.00085317827,0.0021554336,0.0013247215,0.0023161178,0.00051337335],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029211774,0.0001875941,0.0070333336,0.00017313106,0.0001888861,0.0002123894,0.00050740864,0.79551756,0.041721478,0.0056944857,0.001974578,0.14649709],"study_design_scores_gemma":[0.000010020773,0.00006797315,0.0009838765,0.0000099315885,0.000015704325,0.00003909211,0.00006005644,0.9869732,0.0046252124,0.00661093,0.00058764266,0.000016401802],"about_ca_topic_score_codex":0.0027174526,"about_ca_topic_score_gemma":0.0032643853,"teacher_disagreement_score":0.0027174526,"about_ca_system_score_codex":0.00043597622,"about_ca_system_score_gemma":0.0008668119,"threshold_uncertainty_score":0.0068402886},"labels":[],"label_agreement":null},{"id":"W4385571404","doi":"10.18653/v1/2023.acl-long.490","title":"BLIND: Bias Removal With No Demographics","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation","keywords":"Demographics; Computer science; Task (project management); Artificial intelligence; Process (computing); Gender bias; Annotation; Machine learning; Debiasing; Psychology; Social psychology; Demography","score_opus":0.07052993259609026,"score_gpt":0.2664610675017185,"score_spread":0.19593113490562825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571404","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049151048,0.001365472,0.92838466,0.0012884544,0.0007580299,0.00054138165,0.0013924938,0.013963384,0.0031551213],"genre_scores_gemma":[0.4322615,0.0008477313,0.5376182,0.003644947,0.0011796749,0.0011053296,0.0052090795,0.003514657,0.014618883],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949071,0.0017899984,0.00028894123,0.0015171373,0.0009867582,0.00051002303],"domain_scores_gemma":[0.9865023,0.0049485564,0.0009077009,0.004943502,0.002182375,0.0005154697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010273149,0.0030559362,0.002104427,0.002123531,0.0017127026,0.0022586917,0.0032124552,0.0030330976,0.0039449306],"category_scores_gemma":[0.036136705,0.001163505,0.002544506,0.0011826651,0.0017579803,0.005006549,0.0053344034,0.0037097635,0.004990853],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014937724,0.00073247607,0.03396949,0.00078702706,0.0008282123,0.00045349542,0.0021234106,0.04851358,0.029817227,0.01838555,0.07361349,0.78928226],"study_design_scores_gemma":[0.00047039217,0.0006373441,0.010325156,0.00045658718,0.00059626455,0.0010194387,0.0007333942,0.7877276,0.046983466,0.0890501,0.06165977,0.00034044447],"about_ca_topic_score_codex":0.005501724,"about_ca_topic_score_gemma":0.008539327,"teacher_disagreement_score":0.010273149,"about_ca_system_score_codex":0.0009087087,"about_ca_system_score_gemma":0.004185615,"threshold_uncertainty_score":0.05433023},"labels":[],"label_agreement":null},{"id":"W4385571455","doi":"10.18653/v1/2023.findings-acl.150","title":"Attribute Controlled Dialogue Prompting","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Vector Institute","funders":"Vector Institute; University of Waterloo; Government of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Conversation; Task (project management); Domain (mathematical analysis); Open domain; Code (set theory); Artificial intelligence; Control (management); Natural language processing; Language model; Human–computer interaction; Machine learning; Programming language; Question answering; Linguistics","score_opus":0.045212504855818396,"score_gpt":0.2669467549645016,"score_spread":0.2217342501086832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571455","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0316851,0.00055764965,0.9021322,0.00025334672,0.00047315154,0.00053444126,0.0013497028,0.05930533,0.0037091053],"genre_scores_gemma":[0.5106521,0.00031461244,0.46867558,0.0006297078,0.00032746667,0.001230263,0.004895849,0.0041251983,0.009149272],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99670285,0.0013659899,0.00016673714,0.0012575277,0.00034846432,0.00015841718],"domain_scores_gemma":[0.99302965,0.0041727363,0.00030067074,0.0012451824,0.00094027066,0.0003115082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003336655,0.0017592484,0.0011763951,0.000685219,0.0005118557,0.0012533929,0.002008616,0.001237104,0.011078203],"category_scores_gemma":[0.019296482,0.00041442414,0.00077988574,0.00056407927,0.0007018533,0.002075169,0.0021660342,0.0022482644,0.004895549],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020630104,0.00047346877,0.0040436415,0.0010794438,0.00011482584,0.00043655606,0.002526082,0.040871486,0.0696586,0.009474909,0.0417328,0.8275252],"study_design_scores_gemma":[0.00052777614,0.00071649515,0.002433831,0.00014782274,0.00013222342,0.0006185641,0.0007270631,0.80565685,0.07859669,0.034111347,0.07615752,0.0001738402],"about_ca_topic_score_codex":0.00073919387,"about_ca_topic_score_gemma":0.00097411254,"teacher_disagreement_score":0.011078203,"about_ca_system_score_codex":0.0005247076,"about_ca_system_score_gemma":0.00089675514,"threshold_uncertainty_score":0.03706032},"labels":[],"label_agreement":null},{"id":"W4385571607","doi":"10.18653/v1/2023.findings-acl.97","title":"SERENGETI: Massively Multilingual Language Models for Africa","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Advanced Micro Devices","keywords":"Computer science; Set (abstract data type); Natural language processing; Task (project management); Artificial intelligence; Similarity (geometry); Language model; Natural language","score_opus":0.06771763747740082,"score_gpt":0.29756966483051106,"score_spread":0.22985202735311022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571607","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3351111,0.0082642315,0.43703195,0.006717981,0.0037536249,0.0010735185,0.08138852,0.09563801,0.03102113],"genre_scores_gemma":[0.6727493,0.0016377469,0.18357642,0.0015693235,0.0005315384,0.0013179686,0.115126975,0.003919554,0.019571116],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995939,0.00016169244,0.000016896496,0.00014128415,0.000031666346,0.000054539596],"domain_scores_gemma":[0.9993407,0.00034863266,0.000032271615,0.00011034391,0.00010426379,0.00006378857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012252552,0.0023628476,0.0006959978,0.0010155685,0.00086049637,0.0012909792,0.0020334967,0.0010187416,0.007719918],"category_scores_gemma":[0.003447008,0.0006746462,0.0018011522,0.00077079545,0.00037059234,0.0029632614,0.0019014721,0.003373475,0.0044799936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001296516,0.00065654254,0.014843152,0.0008856729,0.0011289369,0.0009570884,0.0012551766,0.29591164,0.00961355,0.010587814,0.31160992,0.35125393],"study_design_scores_gemma":[0.0002548755,0.00019404925,0.0031547311,0.00016214272,0.00015510502,0.00026315326,0.00038265172,0.92244065,0.004864667,0.017242057,0.050794188,0.00009163204],"about_ca_topic_score_codex":0.018470695,"about_ca_topic_score_gemma":0.032920357,"teacher_disagreement_score":0.018470695,"about_ca_system_score_codex":0.0009771205,"about_ca_system_score_gemma":0.0013275987,"threshold_uncertainty_score":0.036726415},"labels":[],"label_agreement":null},{"id":"W4385571655","doi":"10.18653/v1/2023.findings-acl.46","title":"The Web Can Be Your Oyster for Improving Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Renmin University of China; National Natural Science Foundation of China","keywords":"Computer science; Salient; Language model; Task (project management); Information retrieval; ENCODE; Machine learning; Search engine; Artificial intelligence","score_opus":0.05366418091734578,"score_gpt":0.28056988894289203,"score_spread":0.22690570802554624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571655","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07022582,0.008681216,0.8560348,0.0037619842,0.00095650967,0.00021130702,0.0038514873,0.039986867,0.016289983],"genre_scores_gemma":[0.5304942,0.0050001815,0.42936733,0.0020490226,0.0005476808,0.00042267924,0.011363123,0.0024995806,0.018256223],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936396,0.0002072077,0.00004403804,0.00021373422,0.0001259102,0.000045122284],"domain_scores_gemma":[0.9973164,0.0012598064,0.00009942602,0.0008709243,0.00035841684,0.00009495661],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017279628,0.0016436317,0.0007487679,0.0015412633,0.00052347296,0.0017404691,0.001568542,0.001331613,0.009367543],"category_scores_gemma":[0.008435614,0.0006748433,0.0013027001,0.0014175646,0.00049647724,0.0069447225,0.002017867,0.0027899116,0.0067460285],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004236213,0.00044428447,0.004127015,0.00040368934,0.00030172284,0.00023205463,0.00019702298,0.14427863,0.011364959,0.005974155,0.045076758,0.787176],"study_design_scores_gemma":[0.00005271378,0.00016054863,0.0012479316,0.00008345421,0.0000937867,0.00010643999,0.00007613739,0.9525538,0.008593861,0.014842417,0.02214251,0.000046281675],"about_ca_topic_score_codex":0.009467509,"about_ca_topic_score_gemma":0.014599615,"teacher_disagreement_score":0.009467509,"about_ca_system_score_codex":0.0005766524,"about_ca_system_score_gemma":0.0010262433,"threshold_uncertainty_score":0.03133762},"labels":[],"label_agreement":null},{"id":"W4385571753","doi":"10.18653/v1/2023.acl-industry.65","title":"NAG-NER: a Unified Non-Autoregressive Generation Framework for Various NER Tasks","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Decoding methods; Benchmark (surveying); Named-entity recognition; Generative grammar; Sequence (biology); Encoder; Autoregressive model; Generative model; Code (set theory); Field (mathematics); Set (abstract data type); Artificial intelligence; Algorithm; Task (project management); Programming language; Mathematics","score_opus":0.04580949278770992,"score_gpt":0.2975541275448818,"score_spread":0.2517446347571719,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571753","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019598831,0.00032314562,0.9892545,0.00016959505,0.000086347674,0.000059980764,0.0005708178,0.0060524233,0.0015233598],"genre_scores_gemma":[0.12860458,0.00081116194,0.8452011,0.00062602997,0.00029037174,0.00035550093,0.00726241,0.0026355651,0.014213366],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999035,0.00034171273,0.000059355232,0.00033483794,0.00015262092,0.00007644561],"domain_scores_gemma":[0.9986393,0.00072172075,0.00008045171,0.0003016529,0.00020295617,0.000053919517],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024475034,0.0015305971,0.0011917079,0.001713536,0.00070995284,0.001344877,0.0030714439,0.002034335,0.008658995],"category_scores_gemma":[0.0037956778,0.00075011275,0.0026211033,0.001337787,0.00081629073,0.0023979305,0.0020471711,0.0024560986,0.0060995854],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031335116,0.0001893674,0.0015124108,0.00047263294,0.00031623343,0.00071898114,0.0004000548,0.2874858,0.014742592,0.056648705,0.034059808,0.60314006],"study_design_scores_gemma":[0.000026116126,0.000058267644,0.00034963625,0.000031182004,0.00006560649,0.0002497425,0.000036039877,0.94747293,0.006207693,0.033251073,0.012199334,0.000052438358],"about_ca_topic_score_codex":0.0046698456,"about_ca_topic_score_gemma":0.011923226,"teacher_disagreement_score":0.008658995,"about_ca_system_score_codex":0.00071945996,"about_ca_system_score_gemma":0.0010356659,"threshold_uncertainty_score":0.028967261},"labels":[],"label_agreement":null},{"id":"W4385571765","doi":"10.18653/v1/2023.acl-industry.71","title":"Exploring Zero and Few-shot Techniques for Intent Classification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Minnow Environmental (Canada)","funders":"","keywords":"Computer science; Constraint (computer-aided design); Resource (disambiguation); Zero (linguistics); Adaptation (eye); Shot (pellet); Limited resources; Face (sociological concept); Language model; Sample (material); Artificial intelligence; Machine learning; Mathematics; Linguistics; Statistics","score_opus":0.4411939191832756,"score_gpt":0.3375341698411399,"score_spread":0.10365974934213573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571765","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05243711,0.0012034203,0.92950356,0.0005280992,0.00021147038,0.00020274165,0.00049314403,0.013222107,0.0021982614],"genre_scores_gemma":[0.44009793,0.00040065584,0.5452598,0.0007539837,0.00025726896,0.00038214683,0.0045211893,0.0012803411,0.0070467433],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997799,0.0007867292,0.00011502702,0.00075368356,0.000338318,0.00020732624],"domain_scores_gemma":[0.9942504,0.003611272,0.0001731635,0.0010186026,0.00065106375,0.00029544556],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033699258,0.0018945752,0.0015581308,0.0020980092,0.0011850309,0.0016711751,0.0035006048,0.0016752748,0.004200251],"category_scores_gemma":[0.010392423,0.00081109966,0.0014989618,0.0012412977,0.0012532969,0.0052362354,0.0033087647,0.0040779617,0.0029029148],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010144665,0.0011490172,0.004749223,0.0006177799,0.00022797358,0.0002627265,0.0011746294,0.071446,0.023719706,0.011196534,0.021186665,0.86325526],"study_design_scores_gemma":[0.000054406897,0.0001876346,0.0007676404,0.000030196088,0.000043217286,0.0001229583,0.00032730238,0.9689511,0.007977308,0.017972033,0.003521867,0.000044276007],"about_ca_topic_score_codex":0.007999845,"about_ca_topic_score_gemma":0.013972485,"teacher_disagreement_score":0.007999845,"about_ca_system_score_codex":0.0011994484,"about_ca_system_score_gemma":0.0015963323,"threshold_uncertainty_score":0.017822087},"labels":[],"label_agreement":null},{"id":"W4385571773","doi":"10.18653/v1/2023.acl-short.145","title":"Diversity-Aware Coherence Loss for Improving Neural Topic Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of the Fraser Valley; University of British Columbia","funders":"","keywords":"Computer science; Coherence (philosophical gambling strategy); Diversity (politics); Artificial neural network; Artificial intelligence; Mathematics; Statistics","score_opus":0.06884440324724243,"score_gpt":0.25894346780443056,"score_spread":0.1900990645571881,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571773","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.077536635,0.0083607305,0.9048326,0.001156927,0.00033152933,0.000094650524,0.00079096784,0.003407461,0.003488498],"genre_scores_gemma":[0.76880634,0.0024833218,0.21444084,0.00073490734,0.001246992,0.00027596438,0.0045935633,0.0008647515,0.0065534753],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986998,0.0005517187,0.000069277194,0.00031676496,0.00023599697,0.00012649235],"domain_scores_gemma":[0.9958365,0.0027173173,0.00017387829,0.0005233278,0.0005792418,0.00016973153],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00338474,0.0012791314,0.0016516399,0.0017201263,0.000865915,0.0012979301,0.0021232339,0.0014787038,0.0023980902],"category_scores_gemma":[0.011675049,0.0006419986,0.0008454271,0.002067,0.0006229292,0.004132661,0.0024589708,0.002632097,0.0014971944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010349989,0.00034065833,0.0037920047,0.0003502866,0.00044795245,0.00011415544,0.0004933863,0.24451675,0.013299601,0.014030003,0.035733942,0.68584627],"study_design_scores_gemma":[0.000075820615,0.0001281745,0.00063619163,0.000021014603,0.00008798306,0.000047888465,0.000051822655,0.9800827,0.0024230373,0.014300619,0.002127315,0.00001745786],"about_ca_topic_score_codex":0.0048368387,"about_ca_topic_score_gemma":0.0091958875,"teacher_disagreement_score":0.0048368387,"about_ca_system_score_codex":0.0008936057,"about_ca_system_score_gemma":0.0011332995,"threshold_uncertainty_score":0.017900407},"labels":[],"label_agreement":null},{"id":"W4385571841","doi":"10.18653/v1/2023.acl-long.758","title":"Learning New Skills after Deployment: Improving open-domain internet-driven dialogue with human feedback","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Software deployment; Computer science; The Internet; Domain (mathematical analysis); Artificial intelligence; Machine learning; World Wide Web; Software engineering","score_opus":0.01701245785068708,"score_gpt":0.2501860809956955,"score_spread":0.23317362314500842,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571841","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3170451,0.0020790885,0.65786314,0.0017141345,0.00023353295,0.00035417656,0.00076094805,0.014736332,0.005213506],"genre_scores_gemma":[0.8947198,0.00020906034,0.09933842,0.0005662303,0.00011009374,0.00019491547,0.001526525,0.00044143523,0.002893511],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99778867,0.0012871873,0.00006794483,0.00052015076,0.00016015295,0.00017597851],"domain_scores_gemma":[0.9849631,0.011435723,0.000469824,0.0015677141,0.0009847746,0.0005789141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054952702,0.0017280929,0.001381145,0.0008764463,0.0005916317,0.0014001789,0.0021443367,0.0018497425,0.0017632146],"category_scores_gemma":[0.022166274,0.00073020265,0.00079575647,0.00058172474,0.0008847104,0.003765807,0.001959868,0.0030996571,0.0016886998],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014927929,0.0016991956,0.014231532,0.00047745253,0.00031554903,0.00018346761,0.0015524054,0.5547531,0.009248724,0.0028944162,0.012709503,0.40044183],"study_design_scores_gemma":[0.000047338537,0.00014173695,0.00052901974,0.000011576283,0.00002092669,0.000020430793,0.00006563373,0.99537057,0.0014610458,0.0017860676,0.00053264084,0.000012988262],"about_ca_topic_score_codex":0.008711228,"about_ca_topic_score_gemma":0.010432669,"teacher_disagreement_score":0.008711228,"about_ca_system_score_codex":0.001064126,"about_ca_system_score_gemma":0.0013173611,"threshold_uncertainty_score":0.029062092},"labels":[],"label_agreement":null},{"id":"W4385571894","doi":"10.18653/v1/2023.acl-long.68","title":"ThinkSum: Probabilistic reasoning over sets using large language models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Probabilistic logic; Inference; Computer science; Context (archaeology); Set (abstract data type); Artificial intelligence; Natural language processing; Machine learning; Programming language","score_opus":0.04187799106243501,"score_gpt":0.2986571136122419,"score_spread":0.2567791225498069,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571894","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006175755,0.0004295507,0.9599284,0.0007433518,0.000101213955,0.00020976119,0.0020629724,0.028780442,0.0015685259],"genre_scores_gemma":[0.17871273,0.0005038471,0.8062127,0.00096222915,0.00018375291,0.00063462503,0.008112992,0.0017960262,0.0028811402],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99682355,0.0012186894,0.00024612862,0.00077612174,0.0008200541,0.0001155092],"domain_scores_gemma":[0.9909382,0.00673175,0.00032984393,0.0013501514,0.00045621025,0.00019391079],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043273293,0.00257998,0.0013830693,0.0016844086,0.0009669583,0.0041546356,0.005826963,0.002607367,0.0109226825],"category_scores_gemma":[0.019587401,0.0014228041,0.0042070993,0.0013420684,0.0014301145,0.01017563,0.0049357717,0.006149695,0.0033358193],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010037628,0.00046815482,0.0025170206,0.001556598,0.00075135,0.00058951386,0.0007675707,0.39889604,0.0073324186,0.081089266,0.04828217,0.4567461],"study_design_scores_gemma":[0.00009064475,0.000045836703,0.000117679614,0.000030477448,0.000039439652,0.000055698078,0.000039901854,0.9163606,0.0018691578,0.077576526,0.0037445792,0.00002953615],"about_ca_topic_score_codex":0.009483643,"about_ca_topic_score_gemma":0.019172998,"teacher_disagreement_score":0.0109226825,"about_ca_system_score_codex":0.0019904852,"about_ca_system_score_gemma":0.002691901,"threshold_uncertainty_score":0.03653997},"labels":[],"label_agreement":null},{"id":"W4385571931","doi":"10.18653/v1/2023.acl-short.46","title":"MIReAD: Simple Method for Learning High-quality Representations from Scientific Documents","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Categorization; Simple (philosophy); Representation (politics); Transformer; Information retrieval; Class (philosophy); Quality (philosophy); Artificial intelligence; Scientific literature; Language model; Natural language processing","score_opus":0.10473021212332012,"score_gpt":0.4317125563418593,"score_spread":0.32698234421853917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571931","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025384266,0.0028707257,0.92334104,0.0012233662,0.00038693272,0.00035361398,0.0082565285,0.03526843,0.002915117],"genre_scores_gemma":[0.21239886,0.0018868819,0.73511326,0.0009254631,0.00044124838,0.0009884036,0.03855865,0.0018185544,0.007868676],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986608,0.00036503785,0.00013158808,0.00033973422,0.0004097391,0.00009308541],"domain_scores_gemma":[0.9961552,0.0018008739,0.00028573614,0.0009400296,0.00067525625,0.0001429904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030269383,0.0013678559,0.000960578,0.005364683,0.000584617,0.001783662,0.002194548,0.0020221814,0.0052141813],"category_scores_gemma":[0.013614597,0.0005048287,0.0017175071,0.0046329605,0.0005690695,0.0042068656,0.0022008293,0.002652695,0.004141523],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063381455,0.00045444118,0.0053168717,0.0009042165,0.00047790434,0.00020095443,0.000314742,0.0432754,0.016359989,0.010665802,0.08780849,0.83358747],"study_design_scores_gemma":[0.00027868833,0.00034658823,0.0024085953,0.00014183829,0.00019674313,0.00041836593,0.00018187716,0.8844447,0.026189463,0.038931858,0.04634694,0.00011423796],"about_ca_topic_score_codex":0.0031733885,"about_ca_topic_score_gemma":0.0077976827,"teacher_disagreement_score":0.005364683,"about_ca_system_score_codex":0.0010792428,"about_ca_system_score_gemma":0.002394577,"threshold_uncertainty_score":0.01744318},"labels":[],"label_agreement":null},{"id":"W4385571970","doi":"10.18653/v1/2023.acl-short.72","title":"Using contradictions improves question answering systems","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Contradiction; Question answering; Computer science; Context (archaeology); Inference; Logical consequence; Natural language processing; Natural language; Artificial intelligence; Epistemology; Philosophy","score_opus":0.06175751992041341,"score_gpt":0.29644033590898905,"score_spread":0.23468281598857565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571970","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22109124,0.008967307,0.7131218,0.007550837,0.0006747764,0.00052721833,0.0027605423,0.027385166,0.01792114],"genre_scores_gemma":[0.7176386,0.0014219363,0.26971993,0.0013315411,0.00040351306,0.00015303153,0.0063852905,0.0010683924,0.001877783],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97435546,0.011843928,0.0023953256,0.005023122,0.00571627,0.0006658542],"domain_scores_gemma":[0.90380305,0.07347433,0.0033941644,0.011651013,0.0067731943,0.0009043141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024274945,0.001836413,0.0019628867,0.004278426,0.0014677584,0.0051177386,0.0024548967,0.0029398731,0.0038535714],"category_scores_gemma":[0.12131955,0.001087476,0.0025503773,0.0027748712,0.0014109751,0.01219423,0.007423452,0.0032954172,0.0036030551],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001871771,0.0011659091,0.06365358,0.0032335643,0.0012963469,0.0009202071,0.0066865617,0.07276873,0.052373003,0.045710254,0.033147193,0.717173],"study_design_scores_gemma":[0.000291581,0.0008430468,0.015387286,0.00047123883,0.0011013432,0.0013630508,0.0016233765,0.7201125,0.047803335,0.13652202,0.07416985,0.00031128307],"about_ca_topic_score_codex":0.002487863,"about_ca_topic_score_gemma":0.0025036477,"teacher_disagreement_score":0.024274945,"about_ca_system_score_codex":0.0014927854,"about_ca_system_score_gemma":0.0022755207,"threshold_uncertainty_score":0.1283797},"labels":[],"label_agreement":null},{"id":"W4385572006","doi":"10.18653/v1/2023.bionlp-1.46","title":"GRASUM at BioLaySumm Task 1: Background Knowledge Grounding for Readable, Relevant, and Factual Biomedical Lay Summaries","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Automatic summarization; Readability; Computer science; Relevance (law); Task (project management); Ground; Information retrieval; Data science; Engineering","score_opus":0.06062182398145231,"score_gpt":0.3027965851717317,"score_spread":0.2421747611902794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572006","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033219244,0.018904362,0.16645288,0.0109160235,0.0031595281,0.0055448664,0.63465637,0.10462257,0.022524158],"genre_scores_gemma":[0.034772903,0.0021299939,0.20790815,0.001122808,0.00061934174,0.0031982786,0.7401225,0.0027274932,0.0073984857],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937961,0.0021569158,0.0007157979,0.0017457013,0.0012660364,0.0003193495],"domain_scores_gemma":[0.9813048,0.009384742,0.0013370055,0.0028714924,0.0037816542,0.0013202883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008629402,0.0034789962,0.002255099,0.007613175,0.0022272905,0.004706799,0.002983724,0.0037482611,0.031028287],"category_scores_gemma":[0.044712633,0.00073363964,0.0026447931,0.003471192,0.0010030054,0.0052106115,0.006916066,0.0030952105,0.022326607],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001035849,0.00032697423,0.0034504565,0.010078347,0.0004019996,0.00039557528,0.0007609803,0.0061537446,0.010081274,0.0042199846,0.7096988,0.25339603],"study_design_scores_gemma":[0.0010537888,0.0008604215,0.010763961,0.0027545374,0.00061828934,0.00085023354,0.0012565397,0.06685411,0.02241718,0.024589393,0.86768585,0.00029565176],"about_ca_topic_score_codex":0.006396195,"about_ca_topic_score_gemma":0.008791321,"teacher_disagreement_score":0.031028287,"about_ca_system_score_codex":0.0021003964,"about_ca_system_score_gemma":0.00533217,"threshold_uncertainty_score":0.1038},"labels":[],"label_agreement":null},{"id":"W4385572053","doi":"10.18653/v1/2023.findings-acl.369","title":"Shielded Representations: Protecting Sensitive Attributes Through Iterative Gradient-Based Projection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation; Israel Science Foundation; Open Philanthropy Project","keywords":"Computer science; Projection (relational algebra); ENCODE; Task (project management); Representation (politics); Artificial intelligence; Machine learning; Iterative method; Pattern recognition (psychology); Algorithm","score_opus":0.07561534063835071,"score_gpt":0.31807282552317234,"score_spread":0.24245748488482163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572053","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030385662,0.00019414019,0.96674204,0.00031052373,0.000037571066,0.00007068438,0.000064817876,0.0011878139,0.0010068951],"genre_scores_gemma":[0.58885485,0.0003445837,0.40475094,0.00055089797,0.00011660803,0.00033616228,0.0005995988,0.0004831941,0.0039631794],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984451,0.0006351864,0.00006618655,0.0002672841,0.000432299,0.00015398023],"domain_scores_gemma":[0.9962888,0.0016374072,0.00033645076,0.00094874646,0.0006134366,0.00017519544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030221762,0.0019669957,0.0014565998,0.00082622183,0.0007126348,0.0015892249,0.0020858154,0.001721015,0.0017100484],"category_scores_gemma":[0.012890787,0.0006676246,0.0011457959,0.0008396679,0.0017935137,0.0035142542,0.0034707203,0.0035813677,0.0010575987],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042694344,0.0004004867,0.0030960157,0.00020408191,0.00020979888,0.00020228104,0.0008406566,0.36693668,0.022674363,0.032073494,0.008285213,0.56465],"study_design_scores_gemma":[0.000022137292,0.000086695065,0.00022119808,0.000013752018,0.000019818115,0.000050297323,0.000043624204,0.9687006,0.005970068,0.024123454,0.00073133333,0.000017047596],"about_ca_topic_score_codex":0.0029040675,"about_ca_topic_score_gemma":0.0030616564,"teacher_disagreement_score":0.0030221762,"about_ca_system_score_codex":0.000793317,"about_ca_system_score_gemma":0.0017238386,"threshold_uncertainty_score":0.015982985},"labels":[],"label_agreement":null},{"id":"W4385572106","doi":"10.18653/v1/2023.semeval-1.7","title":"BERTastic at SemEval-2023 Task 3: Fine-Tuning Pretrained Multilingual Transformers Does Order Matter?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Transformer; Artificial intelligence; Fine-tuning; SemEval; Deep learning; Language model; Natural language processing; Machine learning; Framing (construction); Task (project management)","score_opus":0.018104909493448913,"score_gpt":0.26280116793827984,"score_spread":0.24469625844483092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572106","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5908489,0.007611724,0.2714753,0.0035211903,0.002062325,0.0007471943,0.021202797,0.07318011,0.029350465],"genre_scores_gemma":[0.7987657,0.00064017536,0.1287503,0.0011993652,0.00027252323,0.0003930691,0.052188296,0.002079614,0.015711011],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987889,0.00036045504,0.000072875155,0.00049605337,0.0001389258,0.00014277737],"domain_scores_gemma":[0.9974427,0.001145859,0.000118205484,0.00069933885,0.00040360226,0.00019030235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003104746,0.002987785,0.0012216344,0.001081606,0.0008034173,0.0017841876,0.0025038074,0.0021053725,0.007985864],"category_scores_gemma":[0.009473803,0.00066832645,0.0013473256,0.0009870215,0.00072502514,0.00424572,0.0018017646,0.004098631,0.005426776],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003173141,0.0014788089,0.015092649,0.0012878507,0.0006890482,0.00061027025,0.00046672052,0.10752315,0.03314274,0.006165302,0.1317791,0.6985913],"study_design_scores_gemma":[0.00049159885,0.0008237801,0.0057775183,0.00014173475,0.00020460023,0.00034294007,0.00036017524,0.9076243,0.04364722,0.01281957,0.027653178,0.000113314716],"about_ca_topic_score_codex":0.010461057,"about_ca_topic_score_gemma":0.023989009,"teacher_disagreement_score":0.010461057,"about_ca_system_score_codex":0.0013098489,"about_ca_system_score_gemma":0.0015417873,"threshold_uncertainty_score":0.026715338},"labels":[],"label_agreement":null},{"id":"W4385572145","doi":"10.18653/v1/2023.codi-1.10","title":"Discourse Information for Document-Level Temporal Dependency Parsing","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; University of Toronto; Vector Institute; Public Health Agency of Canada","funders":"Public Health Agency; Public Health Agency of Canada","keywords":"Parsing; Dependency grammar; Computer science; Sentence; Natural language processing; Dependency (UML); Artificial intelligence; Profiling (computer programming); Dependency graph; Feature (linguistics); Graph; Linguistics; Theoretical computer science; Programming language","score_opus":0.0555614122377747,"score_gpt":0.3160668927874526,"score_spread":0.26050548054967787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572145","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044340335,0.001355461,0.9333374,0.0009309215,0.0002338053,0.00024743655,0.0037849834,0.009790978,0.005978678],"genre_scores_gemma":[0.33862585,0.0008620452,0.6502206,0.0001659383,0.00016119037,0.00019439007,0.0062489393,0.0008961206,0.0026249245],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989309,0.0004964824,0.00008885944,0.00026327203,0.00016782455,0.00005277043],"domain_scores_gemma":[0.99466115,0.003526618,0.00038443148,0.0005491107,0.00078928604,0.00008932901],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021036593,0.00095067924,0.000512391,0.0022237254,0.0006316367,0.0017764618,0.00069120456,0.0006384158,0.0043769535],"category_scores_gemma":[0.012608667,0.0003770415,0.0005535735,0.001973543,0.00033579083,0.004127683,0.0010854484,0.0012443063,0.0021481856],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055399915,0.00015168638,0.0043899296,0.001024093,0.00010667228,0.00034110298,0.0019621535,0.018872125,0.08175176,0.027349383,0.016106986,0.8473901],"study_design_scores_gemma":[0.00006412938,0.00027380415,0.007784509,0.00048051053,0.00045182815,0.000527445,0.0014598189,0.6953192,0.12870549,0.063630484,0.10111728,0.00018548172],"about_ca_topic_score_codex":0.0030116318,"about_ca_topic_score_gemma":0.004887465,"teacher_disagreement_score":0.0043769535,"about_ca_system_score_codex":0.0008109528,"about_ca_system_score_gemma":0.0014712323,"threshold_uncertainty_score":0.014642358},"labels":[],"label_agreement":null},{"id":"W4385572149","doi":"10.18653/v1/2023.nlrse-1.7","title":"Knowledge-Augmented Language Model Prompting for Zero-Shot Knowledge Graph Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Question answering; Computer science; Shot (pellet); Knowledge graph; Graph; Task (project management); Zero (linguistics); Domain knowledge; Natural language processing; Artificial intelligence; Theoretical computer science; Linguistics","score_opus":0.06291405592591279,"score_gpt":0.33947156824000024,"score_spread":0.27655751231408743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572149","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03547035,0.0009949534,0.9001392,0.0007265929,0.00024545114,0.0003460442,0.0013645488,0.057567887,0.0031449057],"genre_scores_gemma":[0.56981957,0.00038708202,0.41677094,0.0011583305,0.00016370603,0.0004618424,0.0049136593,0.001098528,0.005226295],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99791676,0.0009201093,0.00008576774,0.00065888284,0.0002864902,0.00013207254],"domain_scores_gemma":[0.9949216,0.0033876963,0.00016485124,0.0008719233,0.0004403741,0.00021352855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020040437,0.0018552126,0.0010852064,0.00081733207,0.0005565963,0.001306849,0.002506307,0.0024362216,0.009002102],"category_scores_gemma":[0.012408828,0.0005296842,0.00095062307,0.0005452748,0.00095221214,0.005072605,0.0036021378,0.003387297,0.004084971],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017770431,0.0010057791,0.0027197404,0.0019085482,0.00017542363,0.0008627263,0.0026760725,0.07456271,0.06579726,0.0173357,0.04187637,0.78930265],"study_design_scores_gemma":[0.00017358085,0.0004926778,0.0008960264,0.000074575575,0.000096131895,0.0003552353,0.0005605111,0.8875635,0.027959881,0.05992708,0.02181078,0.000090058056],"about_ca_topic_score_codex":0.0036140843,"about_ca_topic_score_gemma":0.006648053,"teacher_disagreement_score":0.009002102,"about_ca_system_score_codex":0.00091365207,"about_ca_system_score_gemma":0.0013798514,"threshold_uncertainty_score":0.030115008},"labels":[],"label_agreement":null},{"id":"W4385572252","doi":"10.18653/v1/2023.acl-short.133","title":"LI-RAGE: Late Interaction Retrieval Augmented Generation with Explicit Signals for Open-Domain Table Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Security token; Relevance (law); Table (database); Domain (mathematical analysis); Representation (politics); Filter (signal processing); Information retrieval; Set (abstract data type); Inference; Labrador Retriever; Data mining; Artificial intelligence; Theoretical computer science; Computer vision; Programming language; Mathematics","score_opus":0.06102899969721623,"score_gpt":0.3158652028916367,"score_spread":0.25483620319442046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572252","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021440277,0.0030466567,0.8652961,0.00086870376,0.0003732986,0.0004551324,0.0071387533,0.09706247,0.0043185856],"genre_scores_gemma":[0.28251398,0.00079166004,0.65996885,0.0013523515,0.00037857005,0.00063874165,0.037851304,0.0030692013,0.013435304],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985656,0.00050436705,0.00006807442,0.00048734905,0.00026172135,0.00011279459],"domain_scores_gemma":[0.9970535,0.00154013,0.00007847582,0.00093670253,0.0002703186,0.00012098818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029257468,0.0017822345,0.00138017,0.0017218867,0.0005884896,0.0024135332,0.004015391,0.0027138742,0.0114156045],"category_scores_gemma":[0.008101984,0.00061209674,0.0018415177,0.0012190726,0.00077389134,0.0049451296,0.0031722577,0.0031478594,0.010191323],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011849462,0.00062641123,0.0036388212,0.0011448989,0.00037533898,0.00039304348,0.0007320864,0.07646965,0.025664136,0.019057035,0.14907083,0.72164273],"study_design_scores_gemma":[0.00019560153,0.00026197502,0.0008034276,0.000043381257,0.00008787668,0.00023243723,0.00014250998,0.92825407,0.013330416,0.02575712,0.030818801,0.000072368224],"about_ca_topic_score_codex":0.005851408,"about_ca_topic_score_gemma":0.0112881,"teacher_disagreement_score":0.0114156045,"about_ca_system_score_codex":0.0010303324,"about_ca_system_score_gemma":0.0014368335,"threshold_uncertainty_score":0.038188994},"labels":[],"label_agreement":null},{"id":"W4385572313","doi":"10.18653/v1/2023.sustainlp-1.22","title":"Small Character Models Match Large Word Models for Autocomplete Under Memory Constraints","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Character (mathematics); Computer science; Simple (philosophy); Word (group theory); Natural language processing; Natural (archaeology); Natural language; Speech recognition; Arithmetic; Artificial intelligence; Linguistics; Mathematics; History; Philosophy; Epistemology; Archaeology","score_opus":0.14364338211369398,"score_gpt":0.28308636775514334,"score_spread":0.13944298564144936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07702219,0.0016385674,0.9000472,0.0016183783,0.00028778386,0.00023587765,0.0038992579,0.0067032883,0.008547418],"genre_scores_gemma":[0.576622,0.0013644827,0.37544006,0.0006988176,0.0005350974,0.000367441,0.023978686,0.0035822294,0.017411118],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99818534,0.0006103474,0.0001740147,0.0005822015,0.00025537875,0.00019281944],"domain_scores_gemma":[0.9874773,0.0087260185,0.00036051738,0.0021982843,0.00091868255,0.00031918497],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018782949,0.0012374045,0.0016667846,0.0013645549,0.0011582435,0.003217695,0.0017373287,0.0018533514,0.012296725],"category_scores_gemma":[0.013210534,0.0011832402,0.001877715,0.0013668844,0.00083716173,0.008094921,0.002558778,0.003029027,0.008367707],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019688085,0.00056429935,0.0046997527,0.0011159326,0.00040213246,0.0012896961,0.0013813651,0.24922396,0.02303179,0.078692816,0.07573289,0.56189656],"study_design_scores_gemma":[0.00005817736,0.00006550558,0.0004889897,0.000044404438,0.000046890727,0.00022101493,0.00018435097,0.91179985,0.0038337398,0.07709706,0.006132387,0.00002768612],"about_ca_topic_score_codex":0.006270898,"about_ca_topic_score_gemma":0.016984873,"teacher_disagreement_score":0.012296725,"about_ca_system_score_codex":0.00089279126,"about_ca_system_score_gemma":0.0016000598,"threshold_uncertainty_score":0.041136682},"labels":[],"label_agreement":null},{"id":"W4385572407","doi":"10.18653/v1/2023.findings-acl.148","title":"Hence, Socrates is mortal: A Benchmark for Natural Language Syllogistic Reasoning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Syllogism; Computer science; Natural language processing; Artificial intelligence; Paraphrase; Natural language understanding; Natural language; Construct (python library); Deductive reasoning; Benchmark (surveying); Programming language; Linguistics","score_opus":0.02139751244422667,"score_gpt":0.2942475370270555,"score_spread":0.2728500245828288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572407","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07268326,0.005490597,0.024970848,0.00235599,0.0014018918,0.0011782432,0.77427256,0.059305068,0.05834163],"genre_scores_gemma":[0.033129748,0.00051742524,0.033084817,0.00050871394,0.00008740115,0.00053606654,0.9274263,0.0013267376,0.0033827177],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945649,0.0014058293,0.0006848799,0.0014489688,0.0014167594,0.0004784991],"domain_scores_gemma":[0.99019486,0.0035217365,0.0004700848,0.0030112236,0.002100964,0.00070115976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003345989,0.0036287713,0.001103354,0.007827924,0.0021845473,0.0040115104,0.005794766,0.0034421827,0.015329074],"category_scores_gemma":[0.021731388,0.0009080998,0.0033982366,0.006879834,0.0013212331,0.007008891,0.0034615495,0.004307416,0.016562356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074306625,0.0010528865,0.009752278,0.0036424072,0.00033554243,0.000369652,0.00047061956,0.014891131,0.0028822897,0.012364058,0.85208863,0.10140743],"study_design_scores_gemma":[0.0008158088,0.0005474666,0.016892597,0.0010211872,0.00024543016,0.0010742852,0.0010755878,0.15966597,0.013157172,0.027066672,0.7782291,0.00020864377],"about_ca_topic_score_codex":0.027865632,"about_ca_topic_score_gemma":0.047849115,"teacher_disagreement_score":0.027865632,"about_ca_system_score_codex":0.0032494287,"about_ca_system_score_gemma":0.0036367571,"threshold_uncertainty_score":0.05540687},"labels":[],"label_agreement":null},{"id":"W4385572414","doi":"10.18653/v1/2023.repl4nlp-1.19","title":"Effectiveness of Data Augmentation for Parameter Efficient Tuning with Limited Data","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; University of Ottawa; Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Fine-tuning; Context (archaeology); Sentence; Language model; Task (project management); Function (biology); Simple (philosophy); Machine learning; Artificial intelligence","score_opus":0.15053650604398205,"score_gpt":0.3475230354588228,"score_spread":0.19698652941484074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572414","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19127457,0.0021852064,0.7863676,0.0013504605,0.00039180063,0.00030922916,0.0010578688,0.011599624,0.005463586],"genre_scores_gemma":[0.75516945,0.00043043087,0.23832273,0.00062421575,0.00011586315,0.0005046958,0.0021272728,0.0007406485,0.0019646992],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970251,0.0015024658,0.00021596719,0.0007248061,0.00037996325,0.00015167738],"domain_scores_gemma":[0.9914703,0.0051684966,0.0002972027,0.002377013,0.0005213116,0.00016558173],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004884036,0.0016713103,0.0010387349,0.00067111506,0.00055912224,0.0015959251,0.0016738909,0.0014291465,0.002361064],"category_scores_gemma":[0.031413175,0.0006323614,0.0008525581,0.0007495646,0.0010809386,0.003734013,0.0024874,0.003238996,0.0015884148],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013433839,0.0008401363,0.009182845,0.00062573235,0.0004142237,0.0003011513,0.00059515104,0.41530815,0.040635545,0.008417155,0.008647171,0.51368934],"study_design_scores_gemma":[0.00010103834,0.00028909833,0.0014957779,0.00007801729,0.00006897764,0.00013902702,0.00014012597,0.95881325,0.021334073,0.01227629,0.005207799,0.000056661167],"about_ca_topic_score_codex":0.0020010066,"about_ca_topic_score_gemma":0.0029969295,"teacher_disagreement_score":0.004884036,"about_ca_system_score_codex":0.00050201075,"about_ca_system_score_gemma":0.001338386,"threshold_uncertainty_score":0.025829554},"labels":[],"label_agreement":null},{"id":"W4385572479","doi":"10.18653/v1/2023.codi-1.9","title":"Entity-based SpanCopy for Abstractive Summarization to Improve the Factual Consistency","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Huawei Technologies","keywords":"Automatic summarization; Computer science; Consistency (knowledge bases); Relevance (law); Information retrieval; Multi-document summarization; Natural language processing; Component (thermodynamics); Artificial intelligence","score_opus":0.031149169354832626,"score_gpt":0.2804351039777526,"score_spread":0.24928593462291998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572479","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0401883,0.0019632508,0.943494,0.00038135596,0.00020951837,0.00034868714,0.0013112869,0.009475042,0.0026285157],"genre_scores_gemma":[0.31446266,0.00085489167,0.6738053,0.0002185837,0.00035972055,0.00034893304,0.0051111095,0.0006986818,0.0041400976],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99834883,0.00043021588,0.0002315327,0.00045229992,0.0004506272,0.00008656821],"domain_scores_gemma":[0.9945253,0.0017760296,0.0006752311,0.0015473244,0.0013124334,0.00016377818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023432428,0.0012336986,0.0010190571,0.0032701045,0.00082725496,0.0016096935,0.0014389778,0.0009922076,0.004926165],"category_scores_gemma":[0.010611746,0.0002807838,0.0006197471,0.0024190075,0.0005208009,0.0037707158,0.002503855,0.0010577777,0.001886259],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005096866,0.00018524338,0.00256851,0.000922065,0.00019663676,0.00038226094,0.001556358,0.010817492,0.0985079,0.011779565,0.014154088,0.85842013],"study_design_scores_gemma":[0.00019691742,0.001079494,0.009875659,0.0002682747,0.0008807839,0.0011890912,0.0013630443,0.55002195,0.2810624,0.040261198,0.11359444,0.00020673664],"about_ca_topic_score_codex":0.0010139733,"about_ca_topic_score_gemma":0.0012165485,"teacher_disagreement_score":0.004926165,"about_ca_system_score_codex":0.00043838102,"about_ca_system_score_gemma":0.00079291797,"threshold_uncertainty_score":0.016479671},"labels":[],"label_agreement":null},{"id":"W4385572482","doi":"10.18653/v1/2023.bionlp-1.30","title":"Evaluation of ChatGPT on Biomedical Tasks: A Zero-Shot Comparison with Fine-Tuned Generative Transformers","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; York University","keywords":"Automatic summarization; Computer science; Transformer; Generative grammar; Benchmark (surveying); Artificial intelligence; Relationship extraction; Domain (mathematical analysis); Language model; Training set; One shot; Machine learning; Natural language processing; Information extraction; Engineering; Mathematics","score_opus":0.11156961665508827,"score_gpt":0.34779937863760385,"score_spread":0.23622976198251558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572482","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44888988,0.021136811,0.41016474,0.0027827222,0.0025125623,0.0015472017,0.009759415,0.07571449,0.027492158],"genre_scores_gemma":[0.8187223,0.0022126897,0.1352032,0.0017257095,0.0003482831,0.0009645355,0.028183803,0.002871855,0.009767715],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974184,0.0010581834,0.00015425173,0.00082796585,0.00036202432,0.00017918547],"domain_scores_gemma":[0.9920116,0.005794298,0.00013376218,0.0009999044,0.0006094717,0.0004509577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048962925,0.0019812814,0.0017338819,0.0016068575,0.0010012053,0.0016793063,0.0035428805,0.002936835,0.005460963],"category_scores_gemma":[0.016388068,0.0006425413,0.0013065801,0.0010510309,0.0013351903,0.0045164917,0.0033306568,0.0032033983,0.0029136231],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004321202,0.0018571253,0.0074969595,0.0034675095,0.001181871,0.0012705238,0.0010531375,0.22850235,0.026480919,0.0063454043,0.055580705,0.66244227],"study_design_scores_gemma":[0.00058260775,0.0017537423,0.0035103757,0.00016164522,0.00035695263,0.0009687196,0.00057660375,0.9441695,0.019185334,0.012202105,0.016407471,0.00012497575],"about_ca_topic_score_codex":0.008843049,"about_ca_topic_score_gemma":0.011760988,"teacher_disagreement_score":0.008843049,"about_ca_system_score_codex":0.0016997189,"about_ca_system_score_gemma":0.0020904029,"threshold_uncertainty_score":0.025894403},"labels":[],"label_agreement":null},{"id":"W4385572574","doi":"10.18653/v1/2023.findings-acl.823","title":"Better Language Models of Code through Self-Improvement","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Automatic summarization; Code (set theory); Benchmark (surveying); Language model; Artificial intelligence; Machine learning; Code generation; Natural language processing; Programming language","score_opus":0.032611129746918384,"score_gpt":0.2701301621696314,"score_spread":0.23751903242271305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572574","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08742947,0.0026327616,0.86338073,0.0014789031,0.00041051622,0.00019144708,0.002564919,0.036191724,0.005719529],"genre_scores_gemma":[0.618499,0.00084609183,0.35420236,0.0011575733,0.0002539158,0.0005173928,0.012844544,0.002863227,0.00881592],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99902856,0.00025726622,0.00004238483,0.00043641674,0.00014332273,0.00009206061],"domain_scores_gemma":[0.9974267,0.0014172436,0.00015590318,0.0004051333,0.0004939013,0.000101060534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011704036,0.0017210903,0.0008582032,0.0011424256,0.00046300833,0.0012765075,0.0017627195,0.0010744194,0.0026864072],"category_scores_gemma":[0.006959349,0.000633104,0.0014097502,0.00080961466,0.00067147677,0.0030290184,0.0013593588,0.003535656,0.0032295834],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040708075,0.0003772501,0.0041545285,0.000498234,0.00021421965,0.00022129937,0.00050372095,0.43209192,0.025098523,0.007857306,0.03388092,0.49469498],"study_design_scores_gemma":[0.00002129086,0.00003783028,0.00019522758,0.000013259545,0.000016940576,0.000025965111,0.000024046665,0.9909065,0.0032606442,0.003470843,0.0020167695,0.000010687834],"about_ca_topic_score_codex":0.010069884,"about_ca_topic_score_gemma":0.019847501,"teacher_disagreement_score":0.010069884,"about_ca_system_score_codex":0.0010919154,"about_ca_system_score_gemma":0.0018157467,"threshold_uncertainty_score":0.020022511},"labels":[],"label_agreement":null},{"id":"W4385572592","doi":"10.18653/v1/2023.bionlp-1.59","title":"KU-DMIS-MSRA at RadSum23: Pre-trained Vision-Language Model for Radiology Report Summarization","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea","keywords":"Automatic summarization; Computer science; Natural language processing; Artificial intelligence; Biomedical text mining; Programming language; Text mining","score_opus":0.023994151892349793,"score_gpt":0.3078320950951783,"score_spread":0.2838379432028285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572592","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064051256,0.014133493,0.5615861,0.0030113906,0.0032955708,0.0018271906,0.06211118,0.2757982,0.014185631],"genre_scores_gemma":[0.2374577,0.0024955827,0.48556048,0.0017858238,0.0009302155,0.002073533,0.23504049,0.004431821,0.030224357],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998403,0.00040432232,0.00015695821,0.0005552454,0.00026956853,0.00021087263],"domain_scores_gemma":[0.9981812,0.0005120441,0.00008066651,0.00047622278,0.00060659053,0.0001431811],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026079079,0.0036018007,0.002346557,0.0033992073,0.0011944188,0.0022802632,0.0046635703,0.0031590417,0.0131829],"category_scores_gemma":[0.005776655,0.0010680509,0.0034181154,0.0020375426,0.0004186554,0.003508678,0.0023783534,0.003571505,0.019555936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000907997,0.00067443476,0.0011425153,0.00057319377,0.00062469137,0.00021989796,0.00012815371,0.018401721,0.014043645,0.001135477,0.20371354,0.7584347],"study_design_scores_gemma":[0.0004838055,0.00085647084,0.003043875,0.00012997429,0.0006303003,0.0003539782,0.00036083985,0.87316394,0.045560807,0.008109093,0.06710091,0.00020606253],"about_ca_topic_score_codex":0.025709271,"about_ca_topic_score_gemma":0.03706447,"teacher_disagreement_score":0.025709271,"about_ca_system_score_codex":0.0016550457,"about_ca_system_score_gemma":0.0036744168,"threshold_uncertainty_score":0.051119268},"labels":[],"label_agreement":null},{"id":"W4385572691","doi":"10.18653/v1/2022.emnlp-main.68","title":"ELMER: A Non-Autoregressive Pre-trained Language Model for Efficient and Effective Text Generation","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Renmin University of China","keywords":"Security token; Computer science; Autoregressive model; Inference; Layer (electronics); Language model; Dependency (UML); Speedup; Artificial intelligence; Permutation (music); Text generation; Token passing; Natural language processing; Speech recognition; Parallel computing; Computer network; Statistics; Mathematics","score_opus":0.011389701396801354,"score_gpt":0.25876018163890674,"score_spread":0.2473704802421054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572691","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009095147,0.00053395965,0.9799612,0.00021550889,0.00010920345,0.000088286746,0.0003481557,0.0087559605,0.0008925366],"genre_scores_gemma":[0.17308001,0.0006748261,0.8117395,0.0005590254,0.00014807994,0.0004557559,0.0031803134,0.001478644,0.008683788],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934036,0.00022503495,0.00003948577,0.00021746811,0.0001124816,0.00006522096],"domain_scores_gemma":[0.998708,0.0007720361,0.000074171665,0.00017886465,0.00021651336,0.000050385508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010742316,0.0012766813,0.00091981475,0.0007545915,0.00035942128,0.0008291107,0.001757888,0.0011572763,0.0039212494],"category_scores_gemma":[0.003232578,0.00061203085,0.000989873,0.00067949085,0.0004438717,0.0021405465,0.00095002324,0.0022764087,0.0038947652],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003424609,0.00023708974,0.00093492796,0.0003663541,0.00013570093,0.0003464779,0.00024529558,0.31162462,0.028196197,0.0068360553,0.015278089,0.6354568],"study_design_scores_gemma":[0.000021702099,0.000063076026,0.00012775777,0.00001066789,0.000016597736,0.000052870077,0.00001790752,0.9850116,0.00875375,0.003051234,0.0028592283,0.000013561543],"about_ca_topic_score_codex":0.0040949965,"about_ca_topic_score_gemma":0.0067669624,"teacher_disagreement_score":0.0040949965,"about_ca_system_score_codex":0.0005896703,"about_ca_system_score_gemma":0.0010961273,"threshold_uncertainty_score":0.01311785},"labels":[],"label_agreement":null},{"id":"W4385572760","doi":"10.18653/v1/2022.emnlp-main.23","title":"Certified Error Control of Candidate Set Pruning for Two-Stage Relevance Ranking","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Pruning; Ranking (information retrieval); Computer science; Relevance (law); Set (abstract data type); Data mining; Domain (mathematical analysis); Error detection and correction; Artificial intelligence; Machine learning; Algorithm; Mathematics","score_opus":0.04368542010145735,"score_gpt":0.2910072447144538,"score_spread":0.24732182461299648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572760","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03592713,0.0009015926,0.9552848,0.00035894883,0.0001732858,0.00035443218,0.00021753406,0.0044153673,0.00236697],"genre_scores_gemma":[0.55825865,0.00025443188,0.43551424,0.00051487586,0.00025897473,0.0006964101,0.0011081696,0.0010246753,0.0023695892],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9730728,0.009627005,0.0023158526,0.0040833415,0.009513223,0.0013877231],"domain_scores_gemma":[0.88462985,0.06882907,0.0068150396,0.023756566,0.014536747,0.0014328315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01942173,0.0016505341,0.0024146456,0.0026050506,0.0015898622,0.0033371812,0.0041988753,0.0026056985,0.00259075],"category_scores_gemma":[0.12160195,0.0008375843,0.0014317068,0.0019005997,0.0028789626,0.00390717,0.0037362599,0.0032083292,0.001703488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025339823,0.000678383,0.015610597,0.00073980365,0.00029235994,0.00048224122,0.0007449038,0.27236006,0.041766293,0.043252178,0.015270578,0.6062686],"study_design_scores_gemma":[0.00022427698,0.0004931947,0.0021609988,0.00009327633,0.00008142009,0.0004913482,0.000077389144,0.94606733,0.02356281,0.022742344,0.0039223866,0.000083173676],"about_ca_topic_score_codex":0.0035177693,"about_ca_topic_score_gemma":0.0044689463,"teacher_disagreement_score":0.01942173,"about_ca_system_score_codex":0.0017124407,"about_ca_system_score_gemma":0.004451521,"threshold_uncertainty_score":0.10271311},"labels":[],"label_agreement":null},{"id":"W4385572790","doi":"10.18653/v1/2022.emnlp-main.101","title":"A Multilingual Perspective Towards the Evaluation of Attribution Methods in Natural Language Inference","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation","keywords":"Computer science; Attribution; Authorship attribution; Natural language processing; Task (project management); Inference; Perspective (graphical); Artificial intelligence; Focus (optics); Natural language; Psychology","score_opus":0.13066287913312408,"score_gpt":0.46496360325847164,"score_spread":0.33430072412534756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572790","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1235008,0.017878065,0.80919373,0.0040216423,0.0011370364,0.00073770824,0.0057893475,0.009219116,0.028522527],"genre_scores_gemma":[0.5179579,0.0029553638,0.46102324,0.00114405,0.000812753,0.000735475,0.01091629,0.002045535,0.0024093871],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.94468415,0.03784222,0.0036568209,0.004656441,0.008436736,0.00072354486],"domain_scores_gemma":[0.87264854,0.09164494,0.0043451083,0.01671755,0.012689732,0.001954129],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.047901027,0.0021808406,0.0013749818,0.0070695905,0.0020623656,0.0068369033,0.0029687735,0.002274483,0.0037428914],"category_scores_gemma":[0.12642325,0.0006339273,0.0016488472,0.0054487246,0.002263246,0.0103269005,0.0080440845,0.004832166,0.0011825293],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003103446,0.0020618797,0.048848834,0.0068100947,0.0042868373,0.0005116903,0.0044322074,0.09610189,0.026139155,0.06398258,0.03076622,0.7129552],"study_design_scores_gemma":[0.00075214205,0.0024929876,0.030394519,0.0019975381,0.0017410979,0.0013001,0.005077514,0.61400265,0.08742875,0.14902847,0.10516186,0.00062238914],"about_ca_topic_score_codex":0.0053261113,"about_ca_topic_score_gemma":0.00634663,"teacher_disagreement_score":0.047901027,"about_ca_system_score_codex":0.0018967097,"about_ca_system_score_gemma":0.0026216542,"threshold_uncertainty_score":0.25332785},"labels":[],"label_agreement":null},{"id":"W4385572802","doi":"10.18653/v1/2022.emnlp-main.205","title":"Revisiting Pre-trained Language Models and their Evaluation for Arabic Natural Language Processing","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Arabic; Natural language processing; Computer science; Artificial intelligence; Linguistics; Natural language; Natural (archaeology); Philosophy; History; Archaeology","score_opus":0.030231393779269772,"score_gpt":0.30621965395222733,"score_spread":0.2759882601729576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572802","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4792488,0.04098956,0.3546336,0.0034145413,0.0063749272,0.0015075277,0.010513731,0.079543196,0.023774125],"genre_scores_gemma":[0.7579374,0.0064819544,0.19050203,0.0010304417,0.00061067543,0.00087117444,0.027738132,0.0020197264,0.012808538],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977635,0.0008897141,0.00026658326,0.0005080444,0.00038792074,0.00018428035],"domain_scores_gemma":[0.9903831,0.0061753863,0.00015137235,0.00085318106,0.00212063,0.00031630837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044630333,0.002456385,0.0013510738,0.0016966276,0.0009963532,0.0026121284,0.0025301315,0.0017519873,0.0058109583],"category_scores_gemma":[0.014644657,0.0008506597,0.0010961242,0.0012462082,0.00052212004,0.0048427363,0.0018220934,0.0039008823,0.0051488453],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023280159,0.0012035515,0.0038904494,0.0011253932,0.0007838355,0.0003470914,0.00041713956,0.087526314,0.012043175,0.0013055917,0.03479651,0.8542331],"study_design_scores_gemma":[0.00020440095,0.0005574706,0.0020373955,0.00013156339,0.00037252606,0.00019915309,0.00036599793,0.97069395,0.015430087,0.002001946,0.007928603,0.0000769894],"about_ca_topic_score_codex":0.026369289,"about_ca_topic_score_gemma":0.026556818,"teacher_disagreement_score":0.026369289,"about_ca_system_score_codex":0.0013859617,"about_ca_system_score_gemma":0.0021572486,"threshold_uncertainty_score":0.052431583},"labels":[],"label_agreement":null},{"id":"W4385572824","doi":"10.18653/v1/2022.emnlp-main.298","title":"MasakhaNER 2.0: Africa-centric Transfer Learning for Named Entity Recognition","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Blessing; Art history; Art; Philosophy; Theology","score_opus":0.04522403817423372,"score_gpt":0.2313573381089553,"score_spread":0.1861332999347216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572824","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016357757,0.0029560653,0.64786303,0.0012286055,0.0008305451,0.0006270527,0.015489056,0.30347645,0.0111715235],"genre_scores_gemma":[0.15964097,0.002676352,0.67461085,0.0011437444,0.0003176971,0.002062161,0.11143317,0.012356858,0.03575818],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9987363,0.00040996043,0.00008616042,0.00038513175,0.00024249866,0.0001399086],"domain_scores_gemma":[0.9989963,0.0003232028,0.00006167951,0.00034654036,0.00016633802,0.0001058886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032237421,0.0022631914,0.0012137977,0.0028401387,0.0010508782,0.0023605644,0.0035105057,0.001789291,0.021673528],"category_scores_gemma":[0.005221547,0.000867814,0.0019397009,0.0022638664,0.00048322565,0.0045630047,0.0053738933,0.0031147138,0.01915806],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014020689,0.00048546505,0.0021518255,0.00062073214,0.0005381552,0.0006136525,0.00033441407,0.028585346,0.00794386,0.011593485,0.26387715,0.68185383],"study_design_scores_gemma":[0.0004291649,0.0004165957,0.0020472575,0.00017988267,0.00021263468,0.0005269257,0.00031794512,0.7569197,0.024843177,0.0425544,0.17138802,0.00016438187],"about_ca_topic_score_codex":0.007894733,"about_ca_topic_score_gemma":0.008419086,"teacher_disagreement_score":0.021673528,"about_ca_system_score_codex":0.00087481854,"about_ca_system_score_gemma":0.00195623,"threshold_uncertainty_score":0.072505176},"labels":[],"label_agreement":null},{"id":"W4385572863","doi":"10.18653/v1/2022.emnlp-demos.42","title":"TextBox 2.0: A Text Generation Library with Pre-trained Language Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Natural language processing; Chen; Artificial intelligence; Natural language; Programming language; Linguistics; Speech recognition; Philosophy","score_opus":0.017463093221928967,"score_gpt":0.21418811049779243,"score_spread":0.19672501727586347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572863","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0023226694,0.0006400512,0.29263347,0.00016884448,0.00040132375,0.0005603757,0.045051575,0.6540189,0.0042028455],"genre_scores_gemma":[0.052200615,0.0013438942,0.5452498,0.00070396287,0.0003844504,0.0050066477,0.26839107,0.101080276,0.025639227],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99914217,0.00019201617,0.00010829977,0.00026558343,0.00022566959,0.00006626264],"domain_scores_gemma":[0.99771285,0.0012081325,0.00013218263,0.00043036556,0.00037997164,0.0001365708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017849783,0.0039121932,0.0016209139,0.0038639824,0.00063154736,0.0017349924,0.003985015,0.0015603297,0.09155204],"category_scores_gemma":[0.0065805917,0.0016518161,0.001860056,0.0021640963,0.0003447952,0.0033852726,0.00277289,0.00176868,0.06672844],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009924041,0.00029690744,0.0010184536,0.0023931873,0.00043208475,0.000722571,0.00032778378,0.0070675416,0.015834717,0.0034189539,0.5811313,0.3863641],"study_design_scores_gemma":[0.0027869227,0.00075357966,0.003189815,0.00045395968,0.0005200779,0.0015064647,0.0003355103,0.39312464,0.09608142,0.038100887,0.4626161,0.0005306648],"about_ca_topic_score_codex":0.0019118895,"about_ca_topic_score_gemma":0.0030169338,"teacher_disagreement_score":0.09155204,"about_ca_system_score_codex":0.0005662687,"about_ca_system_score_gemma":0.0011615737,"threshold_uncertainty_score":0.3062721},"labels":[],"label_agreement":null},{"id":"W4385572901","doi":"10.18653/v1/2022.emnlp-main.418","title":"TemporalWiki: A Lifelong Benchmark for Training and Evaluating Ever-Evolving Language Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Perplexity; Computer science; Benchmark (surveying); Snapshot (computer storage); Machine learning; Lifelong learning; Artificial intelligence; Training set; Adaptability; Language model; Training (meteorology); Database","score_opus":0.09207660360202875,"score_gpt":0.3187674719700207,"score_spread":0.22669086836799196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572901","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50172395,0.015099571,0.23787588,0.0047997846,0.003974248,0.002040063,0.08963805,0.11077402,0.03407449],"genre_scores_gemma":[0.49045712,0.002288163,0.2925434,0.001311996,0.00038863538,0.0022,0.19650233,0.0049738213,0.009334569],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9930438,0.002878839,0.0008725356,0.0016464541,0.0011465217,0.00041188335],"domain_scores_gemma":[0.98213965,0.010469652,0.0006859752,0.0032790084,0.002649206,0.000776584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009088253,0.0042123664,0.0013210821,0.0035110363,0.001510059,0.0029899238,0.0058345404,0.0038050488,0.0045366827],"category_scores_gemma":[0.037925236,0.0010521169,0.0015994171,0.003398789,0.0013193224,0.0069218455,0.0030510966,0.0036992093,0.0041670613],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016918101,0.0027688919,0.022194035,0.004677764,0.0017887591,0.0010765538,0.0010453829,0.2751092,0.011386581,0.006062764,0.24137704,0.43082124],"study_design_scores_gemma":[0.00039314627,0.00094734837,0.007159214,0.00024367584,0.00020015765,0.00050587225,0.0007177939,0.9219673,0.018903863,0.010551023,0.038237862,0.00017280702],"about_ca_topic_score_codex":0.019581813,"about_ca_topic_score_gemma":0.030429037,"teacher_disagreement_score":0.019581813,"about_ca_system_score_codex":0.002186814,"about_ca_system_score_gemma":0.002508282,"threshold_uncertainty_score":0.048063815},"labels":[],"label_agreement":null},{"id":"W4385572930","doi":"10.18653/v1/2022.emnlp-main.151","title":"Generating Information-Seeking Conversations from Unlabeled Documents","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; National Research Foundation of Korea; National Research Foundation","keywords":"Computer science; Benchmark (surveying); Conversation; Context (archaeology); Baseline (sea); Information retrieval; Key (lock); Code (set theory); Resource (disambiguation); Source code; World Wide Web; Programming language","score_opus":0.014237069424310073,"score_gpt":0.22708426834539377,"score_spread":0.2128471989210837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572930","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23594613,0.010559618,0.5375025,0.0038524691,0.001007266,0.003362328,0.1347943,0.05234643,0.020628955],"genre_scores_gemma":[0.32814822,0.00090565317,0.403775,0.00097544503,0.00031863447,0.002403683,0.25585735,0.0009105749,0.006705526],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9954644,0.0022750043,0.0002401084,0.0013465426,0.00048726232,0.0001867726],"domain_scores_gemma":[0.9920225,0.004556956,0.0002879664,0.0015848248,0.0011442541,0.00040345654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003330299,0.0021429833,0.0013915303,0.0027806791,0.0017544241,0.0017434701,0.002877188,0.002518139,0.004593462],"category_scores_gemma":[0.019276114,0.00064798654,0.0016216317,0.0022766222,0.00087436684,0.00472185,0.0031712612,0.0024504447,0.004117652],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030736513,0.00248436,0.0151482215,0.0059802886,0.00067052303,0.00096688326,0.004381173,0.101247,0.04972302,0.023374356,0.25022137,0.5427292],"study_design_scores_gemma":[0.0004790092,0.0006733065,0.0060524293,0.00032262338,0.00021919055,0.0005672359,0.0026789692,0.7889485,0.036304813,0.032811847,0.13074778,0.00019440278],"about_ca_topic_score_codex":0.013472071,"about_ca_topic_score_gemma":0.022899427,"teacher_disagreement_score":0.013472071,"about_ca_system_score_codex":0.0017816713,"about_ca_system_score_gemma":0.0030022527,"threshold_uncertainty_score":0.02678734},"labels":[],"label_agreement":null},{"id":"W4385572948","doi":"10.18653/v1/2022.blackboxnlp-1.8","title":"Post-hoc analysis of Arabic transformer models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Transformer; Natural language processing; Arabic; Artificial intelligence; Vocabulary; Semitic languages; Arabic languages; Linguistics; Speech recognition; Engineering","score_opus":0.02316346400162472,"score_gpt":0.23432408798979446,"score_spread":0.21116062398816973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572948","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5987428,0.0023103827,0.34280393,0.00230136,0.0011336245,0.0005499774,0.006603884,0.019450383,0.026103642],"genre_scores_gemma":[0.92958474,0.00029700442,0.047496922,0.00039432867,0.00011115372,0.0002434779,0.009032194,0.0013684006,0.011471812],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9989221,0.0002879103,0.000071649185,0.00028380155,0.00021506823,0.00021945967],"domain_scores_gemma":[0.9920902,0.004522664,0.00020155619,0.0011079629,0.0018743902,0.00020322227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023481615,0.0016786428,0.00074012607,0.00087170163,0.000607739,0.0016135074,0.0015911616,0.00094474194,0.013385997],"category_scores_gemma":[0.013859757,0.00045049374,0.001210254,0.00057723554,0.0005141988,0.0017172758,0.0013080966,0.0033425083,0.004851875],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019265993,0.0004393865,0.019668637,0.0005683283,0.00078551995,0.00093564304,0.000580248,0.3570418,0.033957046,0.010777256,0.041372154,0.5319474],"study_design_scores_gemma":[0.00004037334,0.00021369557,0.0029145433,0.000034912002,0.00010098953,0.00011814915,0.00026118036,0.96287376,0.022008087,0.0061711064,0.005237862,0.000025387448],"about_ca_topic_score_codex":0.009942003,"about_ca_topic_score_gemma":0.01710713,"teacher_disagreement_score":0.013385997,"about_ca_system_score_codex":0.0015463592,"about_ca_system_score_gemma":0.0018468106,"threshold_uncertainty_score":0.04478067},"labels":[],"label_agreement":null},{"id":"W4385572953","doi":"10.18653/v1/2022.emnlp-main.39","title":"UnifiedSKG: Unifying and Multi-Tasking Structured Knowledge Grounding with Text-to-Text Language Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":222,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Zhàng; Computer science; Philosophy; Natural language processing; Cognitive science; Chen; Artificial intelligence; Linguistics; Humanities; Psychology; History; China","score_opus":0.03643391698202969,"score_gpt":0.2701383438051173,"score_spread":0.23370442682308762,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572953","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011969587,0.002912422,0.9339834,0.0013936324,0.00040577693,0.00071154087,0.008081856,0.038059372,0.0024823898],"genre_scores_gemma":[0.14697047,0.0017704249,0.80770814,0.0008026956,0.00033239025,0.0010411198,0.035624675,0.0015618873,0.0041881227],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99712366,0.0010262831,0.00025181207,0.00097239256,0.000431505,0.00019437735],"domain_scores_gemma":[0.99423873,0.0034305665,0.00028471358,0.0013489128,0.0004269775,0.00027004845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043880683,0.0032502972,0.0024999254,0.0074407766,0.0015663279,0.004105552,0.0040410184,0.003139193,0.007349144],"category_scores_gemma":[0.014286063,0.0013660463,0.004275914,0.005493436,0.001168468,0.011912322,0.006349316,0.0033117395,0.0038672981],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006501044,0.0005520257,0.0029767854,0.0014518718,0.000988417,0.0006011328,0.00080327765,0.052828077,0.0038079873,0.018080648,0.06680198,0.8504577],"study_design_scores_gemma":[0.00022844526,0.00017240205,0.0009418413,0.00023915985,0.00039780186,0.00014801248,0.0004777646,0.82974935,0.0039318893,0.13698652,0.026635002,0.00009179118],"about_ca_topic_score_codex":0.01662404,"about_ca_topic_score_gemma":0.037397433,"teacher_disagreement_score":0.01662404,"about_ca_system_score_codex":0.0017848434,"about_ca_system_score_gemma":0.0040810965,"threshold_uncertainty_score":0.03305453},"labels":[],"label_agreement":null},{"id":"W4385572965","doi":"10.18653/v1/2022.emnlp-main.82","title":"Maieutic Prompting: Logically Consistent Reasoning with Recursive Explanations","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Naval Information Warfare Center Pacific; Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; Allen Institute for Artificial Intelligence","keywords":"Inference; Correctness; Computer science; Robustness (evolution); Commonsense reasoning; Artificial intelligence; Rule of inference; Machine learning; Natural language processing; Algorithm","score_opus":0.030005265001584864,"score_gpt":0.2339525538757103,"score_spread":0.20394728887412542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385572965","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02186242,0.00043775665,0.9123029,0.0011240061,0.00015382157,0.0003976696,0.0013945338,0.05892584,0.0034009975],"genre_scores_gemma":[0.2794943,0.00021600476,0.7108431,0.0009528208,0.000115501316,0.0003332687,0.0029544078,0.0016374558,0.0034531604],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973864,0.001202165,0.00015594819,0.0007261847,0.00040909878,0.00012016441],"domain_scores_gemma":[0.9824656,0.013108743,0.0005958735,0.0024605338,0.001083658,0.00028549728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041498067,0.0015843998,0.0006864965,0.00084929424,0.00056136586,0.0015200018,0.00308826,0.002235481,0.011311536],"category_scores_gemma":[0.03033782,0.0006416188,0.0010646974,0.0005347731,0.0012689383,0.0040273718,0.003001495,0.0038318993,0.0030530104],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014121066,0.0005651209,0.005420959,0.0018622767,0.00017896715,0.00096526614,0.0024706738,0.099010564,0.03683129,0.05901601,0.0563608,0.73590595],"study_design_scores_gemma":[0.00032464575,0.00027957725,0.0006795545,0.00014704387,0.00008112476,0.00038920727,0.00031593765,0.82248974,0.03514232,0.11010781,0.029968962,0.00007413204],"about_ca_topic_score_codex":0.0018900711,"about_ca_topic_score_gemma":0.0047385097,"teacher_disagreement_score":0.011311536,"about_ca_system_score_codex":0.0009673536,"about_ca_system_score_gemma":0.0024573829,"threshold_uncertainty_score":0.037840843},"labels":[],"label_agreement":null},{"id":"W4385573045","doi":"10.18653/v1/2022.emnlp-main.115","title":"Enhancing Self-Consistency and Performance of Pre-Trained Language Models through Natural Language Inference","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Computer science; Consistency (knowledge bases); Inference; Natural language; Natural language processing; Artificial intelligence; Language model; Natural (archaeology); Cognitive science; Psychology; History; Archaeology","score_opus":0.011908467495459494,"score_gpt":0.24899326325836654,"score_spread":0.23708479576290706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573045","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15368165,0.0042547686,0.8010624,0.0016968591,0.0008993337,0.00019423672,0.00091655646,0.031639546,0.0056546004],"genre_scores_gemma":[0.77929664,0.000600205,0.20815645,0.0008941966,0.0004327722,0.000201111,0.0036123989,0.0028908425,0.003915381],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9940469,0.0029075802,0.0003438728,0.0017931738,0.0005649798,0.00034341594],"domain_scores_gemma":[0.96128714,0.028428791,0.000880963,0.00552618,0.0033382352,0.00053863716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012383178,0.001920375,0.0020261295,0.0015131026,0.00095260056,0.0027023454,0.0038785473,0.0024768584,0.002434468],"category_scores_gemma":[0.046038363,0.0015023313,0.0013848633,0.00096685055,0.00090616423,0.0071849185,0.003079286,0.0049603363,0.0030456942],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017603765,0.0009166639,0.012952109,0.0003693069,0.00087632576,0.0002105351,0.000541825,0.287949,0.015597871,0.0037497045,0.019727683,0.65534866],"study_design_scores_gemma":[0.00006838047,0.00007369925,0.00064515934,0.000018147122,0.00007535594,0.00003809076,0.000040973053,0.9880541,0.0058867233,0.00427971,0.0007976716,0.00002195938],"about_ca_topic_score_codex":0.010542845,"about_ca_topic_score_gemma":0.017752374,"teacher_disagreement_score":0.012383178,"about_ca_system_score_codex":0.0012998211,"about_ca_system_score_gemma":0.0019050454,"threshold_uncertainty_score":0.06548929},"labels":[],"label_agreement":null},{"id":"W4385573057","doi":"10.18653/v1/2022.emnlp-main.249","title":"Improving Passage Retrieval with Zero-Shot Question Generation","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Question answering; Ranking (information retrieval); Artificial intelligence; Security token; Task (project management); Domain (mathematical analysis); Zero (linguistics); Shot (pellet); Open domain; Information retrieval; Simple (philosophy); Mathematics","score_opus":0.028390322715001626,"score_gpt":0.23889744985565753,"score_spread":0.2105071271406559,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573057","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04912339,0.0036366102,0.92089427,0.0004658258,0.00034339912,0.00039262243,0.0008024827,0.021091139,0.0032503046],"genre_scores_gemma":[0.4959269,0.0009563164,0.4835509,0.00088874344,0.0004752231,0.0003920581,0.007007589,0.0011701713,0.009632197],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973041,0.001135919,0.00017072835,0.0006522234,0.0005449487,0.00019209771],"domain_scores_gemma":[0.9955707,0.0022030708,0.00019262954,0.0009708221,0.00089567516,0.00016711222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029343644,0.0016961261,0.0019199289,0.0022004428,0.0006706796,0.0014443625,0.0030022257,0.0020806952,0.0041157673],"category_scores_gemma":[0.012739328,0.00040168414,0.0016738159,0.0011503821,0.00079389475,0.003997416,0.001807899,0.0018862858,0.0039265086],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006767331,0.0010106582,0.0036506962,0.0011663387,0.00034293986,0.0005372497,0.00092313025,0.05988723,0.07199642,0.0082434155,0.033845857,0.8177192],"study_design_scores_gemma":[0.00016499998,0.00070507254,0.0018491001,0.00005005636,0.00022422367,0.0006070175,0.00021218747,0.92154366,0.04618233,0.013241138,0.0151151,0.00010511978],"about_ca_topic_score_codex":0.005087726,"about_ca_topic_score_gemma":0.0060516046,"teacher_disagreement_score":0.005087726,"about_ca_system_score_codex":0.00080931775,"about_ca_system_score_gemma":0.0010758594,"threshold_uncertainty_score":0.015518606},"labels":[],"label_agreement":null},{"id":"W4385573112","doi":"10.18653/v1/2022.emnlp-main.69","title":"Multilingual Relation Classification via Efficient and Effective Prompting","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Computer science; Baseline (sea); Task (project management); Relation (database); Margin (machine learning); Artificial intelligence; Natural language processing; Set (abstract data type); Training set; Class (philosophy); Shot (pellet); Translation (biology); One shot; Machine translation; Machine learning; Data mining","score_opus":0.02034332468962246,"score_gpt":0.2612752863144492,"score_spread":0.24093196162482672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573112","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052624434,0.0007076503,0.892308,0.00049592427,0.0003096333,0.00019254796,0.0009784517,0.049619343,0.0027640022],"genre_scores_gemma":[0.42799953,0.00038870645,0.5543435,0.0005734438,0.00021341485,0.0003497908,0.006837852,0.0015063243,0.007787452],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977258,0.00068069313,0.00011058127,0.0010539546,0.0002722239,0.00015674139],"domain_scores_gemma":[0.99670887,0.0014444118,0.00019955981,0.000836321,0.0006025024,0.00020821746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022194257,0.0018744788,0.0013429698,0.001135099,0.0009779996,0.001473012,0.002104278,0.001226972,0.005459562],"category_scores_gemma":[0.009250997,0.0004872095,0.0010655883,0.0015159445,0.00086465356,0.0048551937,0.0033662445,0.0029563112,0.0052742916],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093816954,0.00035145745,0.0037380566,0.00041628754,0.000045084085,0.00035622553,0.0011501464,0.01738173,0.029011194,0.007886633,0.028786004,0.90993893],"study_design_scores_gemma":[0.00023118818,0.00051715644,0.002713783,0.000098162905,0.00009901153,0.00069416774,0.0015075479,0.8372684,0.053234622,0.067277394,0.036236122,0.00012239288],"about_ca_topic_score_codex":0.0019084767,"about_ca_topic_score_gemma":0.0037396848,"teacher_disagreement_score":0.005459562,"about_ca_system_score_codex":0.0007875608,"about_ca_system_score_gemma":0.0023738435,"threshold_uncertainty_score":0.018264115},"labels":[],"label_agreement":null},{"id":"W4385573173","doi":"10.18653/v1/2022.sustainlp-1.7","title":"Data-Efficient Auto-Regressive Document Retrieval for Fact Verification","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Annotation; Information retrieval; Task (project management); Context (archaeology); Question answering; Precision and recall; Sequence (biology); Natural language processing; Document retrieval; Code (set theory); Artificial intelligence; Component (thermodynamics); Programming language","score_opus":0.06656082727731734,"score_gpt":0.3146764219294113,"score_spread":0.24811559465209398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573173","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022265825,0.0004910891,0.9578327,0.00020672595,0.000056289166,0.00013477982,0.0011910005,0.016453385,0.001368185],"genre_scores_gemma":[0.34249625,0.0004271123,0.6419122,0.00015041846,0.00011148236,0.00025331523,0.0066870376,0.0009351633,0.007027032],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99939513,0.00012219914,0.00003796279,0.00023613426,0.00015129911,0.000057191213],"domain_scores_gemma":[0.9980934,0.00083426305,0.00013642368,0.00054980116,0.00032089654,0.00006514841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016568194,0.00066854974,0.0008062407,0.0010716945,0.0004750553,0.0009579148,0.0020899505,0.0008553475,0.00443054],"category_scores_gemma":[0.0054993657,0.0005549446,0.0011072793,0.00089700386,0.00053425704,0.0026646596,0.0011206714,0.0013906937,0.004762524],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006212937,0.0004764105,0.0036492026,0.00044290983,0.00018049407,0.00032519636,0.00034952385,0.2284496,0.05569944,0.019257424,0.035479188,0.6550693],"study_design_scores_gemma":[0.000016649026,0.000030075504,0.0004274273,0.000008265574,0.000014590365,0.00008690112,0.000017892231,0.9824442,0.009648837,0.004612705,0.002678776,0.000013678436],"about_ca_topic_score_codex":0.0103613585,"about_ca_topic_score_gemma":0.018096827,"teacher_disagreement_score":0.0103613585,"about_ca_system_score_codex":0.00080773066,"about_ca_system_score_gemma":0.0013496946,"threshold_uncertainty_score":0.020602107},"labels":[],"label_agreement":null},{"id":"W4385573227","doi":"10.18653/v1/2022.blackboxnlp-1.24","title":"Probing GPT-3’s Linguistic Knowledge on Semantic Tasks","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University","keywords":"Computer science; Sentence; Natural language processing; Negation; Artificial intelligence; Linguistics; Semantic memory; Semantic role labeling; Psychology; Cognition","score_opus":0.032233787751371076,"score_gpt":0.268688338529988,"score_spread":0.23645455077861693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573227","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7072444,0.00038125087,0.22141704,0.0021207377,0.00022981505,0.0009457194,0.0037125186,0.033629674,0.030318825],"genre_scores_gemma":[0.8532448,0.00018579753,0.1328276,0.0008985407,0.00004282122,0.0006661128,0.0063685863,0.002032756,0.003732982],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9951199,0.0019286788,0.0004675297,0.0011256387,0.0010395922,0.0003186905],"domain_scores_gemma":[0.9570244,0.031488,0.0011617725,0.0068163658,0.003050749,0.0004587143],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071517862,0.0014566801,0.00084933994,0.00073369907,0.0004967852,0.0018702487,0.0026096734,0.0015784882,0.0057070274],"category_scores_gemma":[0.053994518,0.0006972283,0.00095738604,0.00061553787,0.0009819083,0.0057459706,0.002929611,0.0032901904,0.0030556528],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018775685,0.0013128119,0.036943275,0.0014417621,0.00027174273,0.0012773043,0.010583167,0.05334785,0.06701346,0.0111982655,0.03578745,0.77894527],"study_design_scores_gemma":[0.00029358352,0.0009571184,0.024604348,0.00026277857,0.00023146336,0.0013527926,0.002295987,0.795462,0.10076244,0.030042972,0.043458816,0.00027571752],"about_ca_topic_score_codex":0.0030626857,"about_ca_topic_score_gemma":0.002442031,"teacher_disagreement_score":0.0071517862,"about_ca_system_score_codex":0.0011878659,"about_ca_system_score_gemma":0.00130333,"threshold_uncertainty_score":0.037822664},"labels":[],"label_agreement":null},{"id":"W4385573233","doi":"10.18653/v1/2022.emnlp-main.97","title":"On the Transformation of Latent Space in Fine-Tuned NLP Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Artificial intelligence; Space (punctuation); Similarity (geometry); Task (project management); Cluster analysis; Transformation (genetics); Function (biology); Probabilistic latent semantic analysis; Natural language processing; Class (philosophy); Machine learning","score_opus":0.040976105222509766,"score_gpt":0.2265543457583512,"score_spread":0.18557824053584143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573233","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17368458,0.0004189948,0.82137096,0.00082227134,0.000038549555,0.00007313642,0.00016053097,0.0007688024,0.0026622012],"genre_scores_gemma":[0.93243366,0.00021416773,0.06441201,0.0001815553,0.000035414076,0.00010486951,0.00028905462,0.00021951563,0.0021096875],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9991636,0.0004228984,0.000031010848,0.00019622948,0.00008315625,0.00010318884],"domain_scores_gemma":[0.99153614,0.00674521,0.0005317678,0.00071378675,0.00028709354,0.00018596477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023607195,0.00079704693,0.0006887429,0.00056993205,0.00055230124,0.0015600749,0.0010123984,0.0012194216,0.002252788],"category_scores_gemma":[0.016371183,0.00060397096,0.00090119644,0.00046389963,0.0018747984,0.0036022319,0.0018826294,0.003076984,0.00037045032],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000109953806,0.0000711244,0.0019071825,0.00005473106,0.000051589137,0.000101344536,0.00037243552,0.9368603,0.0058633583,0.030142706,0.0005518682,0.023913378],"study_design_scores_gemma":[0.0000059933377,0.000027799391,0.00032529235,0.000006548639,0.0000046756354,0.000016295573,0.000019470948,0.9810084,0.00065967976,0.017770305,0.000148977,0.000006524813],"about_ca_topic_score_codex":0.00508595,"about_ca_topic_score_gemma":0.0051587196,"teacher_disagreement_score":0.00508595,"about_ca_system_score_codex":0.0015250434,"about_ca_system_score_gemma":0.0007563467,"threshold_uncertainty_score":0.012484789},"labels":[],"label_agreement":null},{"id":"W4385573248","doi":"10.18653/v1/2022.emnlp-main.41","title":"When Can Transformers Ground and Compose: Insights from Compositional Generalization Benchmarks","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Microsoft Research","keywords":"Transformer; Computer science; Generalization; Grid; Artificial intelligence; Ground; Computation; Theoretical computer science; Machine learning; Algorithm; Mathematics; Electrical engineering; Engineering","score_opus":0.012896747995611264,"score_gpt":0.20776202091603985,"score_spread":0.19486527292042857,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573248","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5291853,0.00080110173,0.43399736,0.003050337,0.00015435435,0.00025531038,0.0017452325,0.0050308704,0.0257801],"genre_scores_gemma":[0.93238324,0.00016927224,0.06305408,0.00038176076,0.000033799475,0.00010766784,0.001511328,0.00045961168,0.0018993543],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9980856,0.00078624254,0.00011948708,0.0005312363,0.00030052228,0.00017695257],"domain_scores_gemma":[0.99020386,0.006356018,0.00041493247,0.0021790736,0.00047740858,0.00036873497],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034022632,0.0012099677,0.0007759465,0.00060221634,0.00068445405,0.002263244,0.0018839425,0.001716832,0.005672374],"category_scores_gemma":[0.026362708,0.00042356626,0.0013537104,0.0005547595,0.0022361842,0.01122672,0.0031812512,0.0032237435,0.0009954813],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001967693,0.0005938306,0.016080549,0.0011873334,0.00035893367,0.0009249987,0.0035549197,0.35957423,0.025842201,0.29715952,0.018937537,0.27381825],"study_design_scores_gemma":[0.00007639576,0.00015226007,0.0014542391,0.000043224773,0.000055726374,0.0001313473,0.00041128084,0.6326574,0.010009792,0.3516619,0.0033139563,0.000032394153],"about_ca_topic_score_codex":0.004932513,"about_ca_topic_score_gemma":0.0075144633,"teacher_disagreement_score":0.005672374,"about_ca_system_score_codex":0.001387158,"about_ca_system_score_gemma":0.0010536998,"threshold_uncertainty_score":0.018975973},"labels":[],"label_agreement":null},{"id":"W4385573315","doi":"10.18653/v1/2022.case-1.12","title":"GGNN@Causal News Corpus 2022:Gated Graph Neural Networks for Causal Event Classification from Social-Political News Articles","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Irish Centre for High-End Computing; University College Cork; Science Foundation Ireland","keywords":"Computer science; Natural language processing; Artificial intelligence; Causality (physics); Event (particle physics); Graph; Task (project management); Artificial neural network; Theoretical computer science","score_opus":0.056178251858021305,"score_gpt":0.2876693020231136,"score_spread":0.23149105016509228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573315","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5023115,0.0076147653,0.3399779,0.0032054787,0.0023449995,0.0010494004,0.06354689,0.055112936,0.02483602],"genre_scores_gemma":[0.7083579,0.0013155339,0.15031816,0.0004845957,0.0004372624,0.00074300874,0.12207594,0.001023229,0.0152444085],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99960655,0.0001660322,0.000024030072,0.000089073845,0.00006875339,0.000045581935],"domain_scores_gemma":[0.99919504,0.00038688013,0.000054418615,0.00018889192,0.00014163225,0.00003309387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008705329,0.0016696483,0.0004921326,0.0021062323,0.00070686586,0.00075322075,0.0013506666,0.0015793752,0.005955621],"category_scores_gemma":[0.0036151905,0.00040777304,0.0007152103,0.0020603784,0.00045715424,0.0016649313,0.0009720854,0.0013043128,0.0019319281],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000775889,0.0010343482,0.006423065,0.0011129413,0.00045451242,0.0010353493,0.00039101014,0.11063322,0.009782299,0.010963553,0.1779381,0.67945576],"study_design_scores_gemma":[0.00020570395,0.00016375646,0.00318656,0.00004658076,0.00010998932,0.00016232178,0.00019197175,0.9470035,0.00815762,0.014037387,0.026686639,0.000048045367],"about_ca_topic_score_codex":0.016999805,"about_ca_topic_score_gemma":0.029850015,"teacher_disagreement_score":0.016999805,"about_ca_system_score_codex":0.0007316672,"about_ca_system_score_gemma":0.00083913724,"threshold_uncertainty_score":0.033801734},"labels":[],"label_agreement":null},{"id":"W4385573317","doi":"10.18653/v1/2022.emnlp-main.37","title":"QRelScore: Better Evaluating Generated Questions with Deeper Understanding of Context-aware Relevance","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; Université de Montréal","funders":"China Knowledge Centre for Engineering Sciences and Technology; Zhejiang University; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research; National Science Foundation","keywords":"Computer science; Relevance (law); Artificial intelligence; Context (archaeology); Adversarial system; Metric (unit); Sentence; Matching (statistics); Rendering (computer graphics); Generative grammar; Natural language processing; Machine learning; Mathematics","score_opus":0.09211376086439695,"score_gpt":0.29648911037518194,"score_spread":0.204375349510785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573317","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28055963,0.008091433,0.6590029,0.0008381941,0.00068444485,0.001791797,0.0057450505,0.027270632,0.01601593],"genre_scores_gemma":[0.734381,0.0005228852,0.25109026,0.00042258765,0.00018225519,0.00059826917,0.0080581615,0.00091418414,0.0038304473],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99271154,0.0034320736,0.00049654505,0.0014095178,0.0017274015,0.00022292933],"domain_scores_gemma":[0.97368807,0.019789137,0.001440784,0.002196443,0.0022294093,0.0006561506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007522694,0.0019368044,0.0010285532,0.0034839723,0.00048455826,0.0021202224,0.0012048843,0.0020089832,0.0047475873],"category_scores_gemma":[0.051933262,0.00023634889,0.00072906754,0.0010817116,0.00069073366,0.0032559666,0.0023035412,0.0013928832,0.0015634751],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00225809,0.0007742704,0.026996918,0.003130778,0.0006083476,0.00042212554,0.0015074284,0.086462274,0.053108025,0.009669168,0.03260479,0.7824577],"study_design_scores_gemma":[0.00029016813,0.0017521204,0.025226096,0.00023702037,0.00027350173,0.0007785687,0.0005003193,0.8794791,0.049476843,0.01895604,0.022781923,0.00024831042],"about_ca_topic_score_codex":0.0020863614,"about_ca_topic_score_gemma":0.0034897963,"teacher_disagreement_score":0.007522694,"about_ca_system_score_codex":0.0009017195,"about_ca_system_score_gemma":0.0011215375,"threshold_uncertainty_score":0.039784253},"labels":[],"label_agreement":null},{"id":"W4385573343","doi":"10.18653/v1/2022.findings-emnlp.276","title":"Keep Me Updated! Memory Management in Long-term Conversations","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Term (time); Computer science; Linguistics; Physics; Philosophy; Astronomy","score_opus":0.019436825075837334,"score_gpt":0.24677092729368805,"score_spread":0.22733410221785072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573343","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5327774,0.021342568,0.09914303,0.05892653,0.0077108806,0.00046820732,0.007519745,0.018176084,0.2539356],"genre_scores_gemma":[0.9298713,0.0021393257,0.017435988,0.002902517,0.0011269698,0.00020523714,0.002298068,0.001434225,0.04258628],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998741,0.00064406666,0.000078340636,0.00017743865,0.00022339063,0.00013584207],"domain_scores_gemma":[0.9915461,0.0035494387,0.00085346884,0.0015400888,0.0016348313,0.00087610877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020680833,0.00037939902,0.00034080035,0.00087628973,0.0022565501,0.0033877012,0.0010028683,0.0013455368,0.01943747],"category_scores_gemma":[0.026094638,0.00040637315,0.00025437045,0.00090310036,0.0005144022,0.008272439,0.0028573885,0.0010768777,0.011679782],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001488229,0.0001517348,0.028230332,0.00072224473,0.00011416347,0.0014022869,0.030211259,0.0004848966,0.011098435,0.012455646,0.25708213,0.6565587],"study_design_scores_gemma":[0.00017348812,0.00046520948,0.06798475,0.0012574175,0.0005153841,0.0036918845,0.041768298,0.016442861,0.016827852,0.060200676,0.79031223,0.00035994948],"about_ca_topic_score_codex":0.002406636,"about_ca_topic_score_gemma":0.0028297673,"teacher_disagreement_score":0.01943747,"about_ca_system_score_codex":0.0005363625,"about_ca_system_score_gemma":0.0006366058,"threshold_uncertainty_score":0.06502479},"labels":[],"label_agreement":null},{"id":"W4385573345","doi":"10.18653/v1/2022.nllp-1.30","title":"Computing and Exploiting Document Structure to Improve Unsupervised Extractive Summarization of Legal Case Decisions","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"University of Pittsburgh; National Science Foundation","keywords":"Automatic summarization; Computer science; Argumentative; Exploit; Representation (politics); Legal document; Artificial intelligence; Information retrieval; Ranking (information retrieval); Graph; Legal case; Domain (mathematical analysis); Data mining; Machine learning; Natural language processing; Theoretical computer science; Mathematics","score_opus":0.016993625204637198,"score_gpt":0.2676518000611616,"score_spread":0.2506581748565244,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573345","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17156069,0.0089365225,0.7752083,0.0014862267,0.0006809985,0.0010064755,0.016054578,0.017382197,0.007684096],"genre_scores_gemma":[0.372961,0.0026257618,0.5389184,0.00030985937,0.0009738937,0.00064989337,0.07438777,0.0008124714,0.008360996],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988663,0.00024394196,0.00013904141,0.00034670057,0.00029490056,0.000109078246],"domain_scores_gemma":[0.99595225,0.0017872349,0.00046437632,0.0005052388,0.0011664192,0.000124549],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013644685,0.0016313313,0.0013161772,0.01027875,0.0009032187,0.0020782214,0.0012088007,0.0010265885,0.0016356177],"category_scores_gemma":[0.007887711,0.0004144931,0.0010731908,0.0057728416,0.00036566547,0.0030379072,0.00092572265,0.0015334514,0.002211644],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000355195,0.00039916683,0.0070847706,0.00084981124,0.0003305081,0.000302008,0.0007696355,0.025341226,0.025496911,0.0052407207,0.049143147,0.8846869],"study_design_scores_gemma":[0.00022458078,0.0006403534,0.016634226,0.00025656295,0.00079318776,0.0005897277,0.0012933306,0.85372084,0.03484369,0.037207678,0.05365552,0.00014031728],"about_ca_topic_score_codex":0.009066301,"about_ca_topic_score_gemma":0.02376994,"teacher_disagreement_score":0.01027875,"about_ca_system_score_codex":0.0008266282,"about_ca_system_score_gemma":0.0018128089,"threshold_uncertainty_score":0.018027067},"labels":[],"label_agreement":null},{"id":"W4385573352","doi":"10.18653/v1/2022.nllp-1.24","title":"Detecting Relevant Differences Between Similar Legal Texts","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Metadata; Computer science; Task (project management); Terminology; Natural language processing; Sentence; Artificial intelligence; Information retrieval; Focus (optics); Variation (astronomy); Resource (disambiguation); Legal case; Data science; World Wide Web; Linguistics; Political science; Law","score_opus":0.03539416592959268,"score_gpt":0.2492053513097113,"score_spread":0.21381118538011862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573352","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77722335,0.00295928,0.19190827,0.0021767206,0.00054568256,0.00045410168,0.008246349,0.004448977,0.012037266],"genre_scores_gemma":[0.87001574,0.0003816759,0.1113946,0.00036897513,0.00031275296,0.00021573604,0.014643906,0.00031041677,0.0023561767],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9978319,0.00045119846,0.00021421383,0.00087319216,0.00043582957,0.00019365824],"domain_scores_gemma":[0.99291104,0.003802097,0.0011534528,0.0005581061,0.00119393,0.0003813708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017192035,0.0005972636,0.0005720232,0.004404023,0.0010368406,0.0013901878,0.0009462937,0.0015325439,0.0024226292],"category_scores_gemma":[0.009460577,0.00029942903,0.0006935327,0.0015664778,0.000789928,0.0033026985,0.0014864088,0.0013455162,0.0012368108],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021386063,0.00090592203,0.117610894,0.0019614652,0.0003366427,0.003059733,0.005815369,0.010596131,0.16912173,0.02123176,0.0506252,0.6165966],"study_design_scores_gemma":[0.000291425,0.00071093946,0.22299455,0.0005257394,0.0005139409,0.003833635,0.0073094056,0.48692492,0.117919646,0.063121274,0.09558546,0.00026922335],"about_ca_topic_score_codex":0.0033964333,"about_ca_topic_score_gemma":0.0062581007,"teacher_disagreement_score":0.004404023,"about_ca_system_score_codex":0.0011281021,"about_ca_system_score_gemma":0.00096405426,"threshold_uncertainty_score":0.009092152},"labels":[],"label_agreement":null},{"id":"W4385573356","doi":"10.18653/v1/2022.emnlp-main.673","title":"Label-aware Multi-level Contrastive Learning for Cross-lingual Spoken Language Understanding","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"People's Government of Jilin Province; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Utterance; Benchmark (surveying); Natural language processing; Artificial intelligence; Context (archaeology); Word (group theory); Spoken language; Semantics (computer science); Code (set theory); Linguistics; Programming language","score_opus":0.15914910239413696,"score_gpt":0.3535188660146792,"score_spread":0.19436976362054223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573356","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053996757,0.0012659625,0.9363614,0.00029622548,0.0001114651,0.00012407964,0.0004258804,0.0055117644,0.0019065714],"genre_scores_gemma":[0.6135143,0.00043747705,0.37475032,0.0006704633,0.00015260634,0.00037172553,0.0038713086,0.0008241989,0.0054075564],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9975454,0.0008210953,0.0000969835,0.0010408165,0.0003088564,0.00018681085],"domain_scores_gemma":[0.996256,0.0022448145,0.00022011712,0.00060136546,0.00053300103,0.0001447183],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029525822,0.0016457511,0.0015256766,0.0015757768,0.00092677545,0.0015398059,0.0031782975,0.0018661434,0.0020123818],"category_scores_gemma":[0.0074040866,0.0005970912,0.0014325448,0.0012162264,0.0011845987,0.0034360169,0.004281895,0.0039222646,0.0014870415],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006485107,0.00071788655,0.004142117,0.00040257323,0.00033063273,0.00035479447,0.001216809,0.14531605,0.03539112,0.009977096,0.008917851,0.79258466],"study_design_scores_gemma":[0.000024783589,0.00009977039,0.000559615,0.000017540291,0.000035276487,0.000058440706,0.00012848247,0.97903675,0.0074481713,0.01078875,0.0017779041,0.000024618737],"about_ca_topic_score_codex":0.0070532584,"about_ca_topic_score_gemma":0.011720722,"teacher_disagreement_score":0.0070532584,"about_ca_system_score_codex":0.001319587,"about_ca_system_score_gemma":0.0012760577,"threshold_uncertainty_score":0.015614927},"labels":[],"label_agreement":null},{"id":"W4385573402","doi":"10.18653/v1/2022.findings-emnlp.384","title":"XRICL: Cross-lingual Retrieval-Augmented In-Context Learning for Cross-lingual Text-to-SQL Semantic Parsing","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Parsing; SQL; Information retrieval; Programming language","score_opus":0.03103624341915619,"score_gpt":0.32614604897143906,"score_spread":0.2951098055522829,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573402","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06984593,0.00522596,0.67710143,0.0018775528,0.0010471016,0.0013718077,0.01403731,0.2152408,0.014252115],"genre_scores_gemma":[0.32037616,0.0012462704,0.5780222,0.0029906223,0.00033124033,0.0015050571,0.07341438,0.0044723246,0.017641759],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9968636,0.0009634094,0.00016200932,0.001407246,0.00035940768,0.0002443831],"domain_scores_gemma":[0.99650025,0.0013615496,0.00013801428,0.001171121,0.00065186346,0.00017718883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039904034,0.003325739,0.0018879125,0.002412943,0.0013457065,0.0023756758,0.004916045,0.0029230106,0.011059605],"category_scores_gemma":[0.010202522,0.0009773565,0.0023623437,0.0018288672,0.0015237202,0.0068069384,0.006972526,0.0049693887,0.010186198],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074946706,0.0015150886,0.005440701,0.0014050987,0.0005047962,0.0009043993,0.0011926043,0.07373755,0.018772775,0.009880666,0.16947146,0.7164255],"study_design_scores_gemma":[0.0002339641,0.00047843112,0.0017235737,0.000120430035,0.00015977076,0.00042258517,0.0006071138,0.92605734,0.01561512,0.01755214,0.03687579,0.00015379196],"about_ca_topic_score_codex":0.018091286,"about_ca_topic_score_gemma":0.028296633,"teacher_disagreement_score":0.018091286,"about_ca_system_score_codex":0.0020893617,"about_ca_system_score_gemma":0.0032858786,"threshold_uncertainty_score":0.036998093},"labels":[],"label_agreement":null},{"id":"W4385573416","doi":"10.18653/v1/2022.findings-emnlp.393","title":"Detecting Languages Unintelligible to Multilingual Models through Local Structure Probes","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; Huawei Technologies (Canada); Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Set (abstract data type); Variety (cybernetics); Task (project management); Construct (python library); Sentence; Transfer (computing); Transfer of learning; Language model; Programming language","score_opus":0.029890264488388225,"score_gpt":0.2909275560061457,"score_spread":0.2610372915177575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573416","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5245926,0.0010026183,0.45220712,0.0015794323,0.00017274765,0.00016750688,0.0015553875,0.0073870583,0.011335535],"genre_scores_gemma":[0.9360279,0.00018909432,0.056820337,0.0005104651,0.00005994657,0.000118169024,0.0027656727,0.0007029658,0.0028055082],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983438,0.00065537664,0.00006455449,0.0006523492,0.00013687616,0.00014709351],"domain_scores_gemma":[0.9953366,0.0027322033,0.00033400298,0.00088017114,0.00049952697,0.00021753128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002368634,0.0012150997,0.00072321435,0.0010967238,0.001253461,0.0028772193,0.0013459751,0.0012521382,0.0028071401],"category_scores_gemma":[0.008933166,0.00055982376,0.0012855632,0.00092173915,0.001103394,0.004221437,0.0029835792,0.0028147334,0.0018797608],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002008412,0.00082584197,0.0696718,0.0009846617,0.0010717472,0.0025354375,0.008699149,0.13481487,0.13890053,0.04089993,0.03117905,0.5684085],"study_design_scores_gemma":[0.000072711766,0.00031209644,0.009518404,0.00011542559,0.00028616455,0.0006569007,0.003274491,0.8745674,0.03584353,0.060129683,0.0150888525,0.00013434318],"about_ca_topic_score_codex":0.004417529,"about_ca_topic_score_gemma":0.011531905,"teacher_disagreement_score":0.004417529,"about_ca_system_score_codex":0.0011826132,"about_ca_system_score_gemma":0.0013134399,"threshold_uncertainty_score":0.012526631},"labels":[],"label_agreement":null},{"id":"W4385573423","doi":"10.18653/v1/2022.findings-emnlp.73","title":"MCP: Self-supervised Pre-training for Personalized Chatbots with Multi-level Contrastive Sampling","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Fundamental Research Funds for the Central Universities; Ministry of Education, India; Renmin University of China; National Natural Science Foundation of China","keywords":"Computer science; Leverage (statistics); Utterance; Artificial intelligence; Dialog box; Encoder; Focus (optics); Natural language processing; Machine learning; Speech recognition; World Wide Web","score_opus":0.13591573542905944,"score_gpt":0.3056195345867609,"score_spread":0.16970379915770145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040651407,0.00041493546,0.9498843,0.00017410683,0.00008265023,0.00022106666,0.00021622035,0.007373458,0.000981838],"genre_scores_gemma":[0.6076283,0.00018142416,0.38186908,0.0005964313,0.00013535029,0.0009401814,0.002162678,0.0005756044,0.0059109074],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989103,0.00044178005,0.000041167612,0.0003862833,0.00012212739,0.00009819961],"domain_scores_gemma":[0.99738175,0.0015628507,0.00013926468,0.00038929633,0.0003561971,0.00017059714],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021701201,0.0014926305,0.0010765991,0.00066541147,0.0005300564,0.0005831676,0.0023198933,0.0012096755,0.002544069],"category_scores_gemma":[0.0055242456,0.0006646313,0.00088737404,0.0003970225,0.0007540937,0.0013630729,0.0014745418,0.0024560133,0.0014900622],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001105937,0.0010491677,0.005888756,0.00047822268,0.00023060941,0.00024848938,0.0009431022,0.27278855,0.034272116,0.006182603,0.014401163,0.6624113],"study_design_scores_gemma":[0.000023036433,0.0001056456,0.00034549463,0.000008212021,0.000012320205,0.000027274773,0.000029442925,0.9934848,0.003732285,0.0014896344,0.0007331878,0.000008779431],"about_ca_topic_score_codex":0.0032623082,"about_ca_topic_score_gemma":0.006371369,"teacher_disagreement_score":0.0032623082,"about_ca_system_score_codex":0.00076773856,"about_ca_system_score_gemma":0.0011936841,"threshold_uncertainty_score":0.011476815},"labels":[],"label_agreement":null},{"id":"W4385573444","doi":"10.18653/v1/2022.emnlp-main.793","title":"Predicting Fine-Tuning Performance with Probing","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Canadian Institute for Advanced Research","keywords":"Computer science; Artificial intelligence; Language model; Fine-tuning; Proxy (statistics); Machine learning; Deep learning; Natural language processing","score_opus":0.017868475209527685,"score_gpt":0.20189179532455784,"score_spread":0.18402332011503014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573444","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8222804,0.003035613,0.14741255,0.0011549601,0.00032329487,0.00026530295,0.0030349027,0.013979604,0.008513373],"genre_scores_gemma":[0.96986747,0.00023892488,0.024522144,0.00021918332,0.000052058465,0.00013851048,0.0030871565,0.00057262543,0.001302021],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99679655,0.0011190339,0.0002053406,0.0009924317,0.0004997618,0.0003868222],"domain_scores_gemma":[0.9747663,0.01639539,0.0011745837,0.005093942,0.0019315735,0.00063818664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058750473,0.0017090816,0.0009058607,0.0013420627,0.00050906284,0.0018059657,0.0011724662,0.0021057695,0.0014008278],"category_scores_gemma":[0.041990887,0.00057332806,0.000650931,0.0012301772,0.00096034555,0.0042785397,0.0018013143,0.002991448,0.0022923574],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015910073,0.00063725683,0.07891227,0.00052199373,0.00046796992,0.00033421456,0.00094323687,0.52434576,0.036492724,0.0040645543,0.019977812,0.33171123],"study_design_scores_gemma":[0.000049789855,0.00049777154,0.013933562,0.000047555415,0.00008082061,0.00017118064,0.0001881603,0.94975656,0.022347478,0.00908853,0.0037684906,0.000070060705],"about_ca_topic_score_codex":0.0039492687,"about_ca_topic_score_gemma":0.0036634528,"teacher_disagreement_score":0.0058750473,"about_ca_system_score_codex":0.0008802397,"about_ca_system_score_gemma":0.0007466328,"threshold_uncertainty_score":0.03107059},"labels":[],"label_agreement":null},{"id":"W4385573473","doi":"10.18653/v1/2022.emnlp-main.701","title":"Saving Dense Retriever from Shortcut Dependency in Conversational Search","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Dependency (UML); Computer science; Exploit; Recall; Artificial intelligence; Cognitive psychology; Computer security; Psychology","score_opus":0.0343037280176667,"score_gpt":0.2557759169970624,"score_spread":0.22147218897939572,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573473","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37849635,0.009421779,0.52703017,0.0025871277,0.00030982486,0.0007142071,0.005199123,0.060868252,0.01537319],"genre_scores_gemma":[0.8514157,0.00092904817,0.12703341,0.0010884366,0.00020170493,0.00029417974,0.009313356,0.0015589886,0.008165191],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99807644,0.00072356826,0.00014307695,0.00048322653,0.00040345217,0.00017017953],"domain_scores_gemma":[0.99203616,0.0052315304,0.00026969545,0.0014801275,0.00075215404,0.00023027942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002789283,0.0021838902,0.0017228576,0.0016386228,0.00081056345,0.0015294972,0.002459035,0.0020517411,0.004793477],"category_scores_gemma":[0.014017644,0.0007426105,0.0011485252,0.001024407,0.00085599197,0.0075790132,0.0029102238,0.0023455792,0.003784776],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022210912,0.0008763771,0.010669267,0.0019261836,0.00039937758,0.0012316129,0.0020580974,0.0973038,0.045795158,0.008989039,0.045070503,0.78345954],"study_design_scores_gemma":[0.0001476545,0.0006378191,0.002152229,0.00008327401,0.00018544272,0.0008247532,0.0004718939,0.94889385,0.017405292,0.017779393,0.01131979,0.00009860837],"about_ca_topic_score_codex":0.00875401,"about_ca_topic_score_gemma":0.012901968,"teacher_disagreement_score":0.00875401,"about_ca_system_score_codex":0.00087786803,"about_ca_system_score_gemma":0.0017530866,"threshold_uncertainty_score":0.017406106},"labels":[],"label_agreement":null},{"id":"W4385573486","doi":"10.18653/v1/2022.findings-emnlp.191","title":"Partially-Random Initialization: A Smoking Gun for Binarization Hypothesis of BERT","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Initialization; Embedding; Encoder; Task (project management); Set (abstract data type); Artificial intelligence; Enhanced Data Rates for GSM Evolution; Binary number; Edge device; Machine learning; Programming language; Operating system","score_opus":0.04659736887388309,"score_gpt":0.2515191874708972,"score_spread":0.20492181859701408,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573486","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11508777,0.0017853217,0.8660294,0.0017499704,0.00037735046,0.00014931355,0.00056686683,0.0074456874,0.0068083866],"genre_scores_gemma":[0.82558715,0.0004361598,0.16231343,0.001362752,0.00017253509,0.00023647565,0.0018983051,0.0010012517,0.006991871],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99915385,0.00025455415,0.00003643718,0.00030670647,0.0001225379,0.00012585541],"domain_scores_gemma":[0.9978034,0.00093407393,0.0001635964,0.0006913077,0.00027545134,0.00013220232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014926078,0.0015697432,0.0010201888,0.0005128603,0.0006365424,0.0011835633,0.001960006,0.0014673462,0.00331542],"category_scores_gemma":[0.0091970125,0.00060159474,0.0006249588,0.0004955388,0.0014870182,0.0033247198,0.0017111404,0.003047503,0.0017880508],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002294004,0.00041011994,0.0059234947,0.00041992794,0.00024899514,0.00043245457,0.0004086239,0.55889547,0.030069219,0.06532047,0.033035412,0.30254185],"study_design_scores_gemma":[0.00004540955,0.000109424545,0.00047091668,0.00003176504,0.000021556034,0.00008725516,0.000023348297,0.9745044,0.007947485,0.014927002,0.0018098269,0.00002156323],"about_ca_topic_score_codex":0.003936312,"about_ca_topic_score_gemma":0.006282607,"teacher_disagreement_score":0.003936312,"about_ca_system_score_codex":0.00095998566,"about_ca_system_score_gemma":0.0012241831,"threshold_uncertainty_score":0.011091173},"labels":[],"label_agreement":null},{"id":"W4385573538","doi":"10.18653/v1/2022.emnlp-main.619","title":"Intriguing Properties of Compression on Multilingual Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Compression (physics); Robustness (evolution); Generalization; Multilingualism; Language model; Data compression; Data compression ratio; Scaling; Natural language processing; Artificial intelligence; Image compression; Linguistics; Mathematics","score_opus":0.07661469801205291,"score_gpt":0.25831594561231275,"score_spread":0.18170124760025985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573538","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5132662,0.0014400568,0.45868754,0.0049120053,0.00022265941,0.00018307849,0.001745793,0.0031100009,0.016432635],"genre_scores_gemma":[0.94348764,0.0005647842,0.05113087,0.00045032694,0.00017121587,0.00013321418,0.0017931296,0.00030385563,0.0019650366],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99888164,0.00033328103,0.000090922826,0.00029534978,0.0002838014,0.00011505083],"domain_scores_gemma":[0.9882839,0.006957948,0.000657242,0.0032374365,0.00066672097,0.0001967509],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028854362,0.0007006958,0.00059841695,0.00067425123,0.0006514197,0.0010862015,0.0007765963,0.0007225134,0.003087199],"category_scores_gemma":[0.021607338,0.00033917034,0.0005094757,0.0008981344,0.0015673558,0.003216935,0.0017629889,0.0020366856,0.0007167872],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061559794,0.00022880222,0.021919731,0.00040030765,0.00024920853,0.0012760535,0.001628279,0.5559548,0.052807976,0.07770382,0.01153157,0.27568385],"study_design_scores_gemma":[0.000038858834,0.00017228126,0.008572039,0.00006179648,0.00005297037,0.00070585014,0.00042874427,0.87257636,0.021542853,0.089920975,0.005864578,0.00006272356],"about_ca_topic_score_codex":0.0029604435,"about_ca_topic_score_gemma":0.0041234926,"teacher_disagreement_score":0.003087199,"about_ca_system_score_codex":0.0005533749,"about_ca_system_score_gemma":0.00076919986,"threshold_uncertainty_score":0.015259802},"labels":[],"label_agreement":null},{"id":"W4385573576","doi":"10.18653/v1/2022.findings-emnlp.146","title":"MatRank: Text Re-ranking by Latent Preference Matrix","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Ranking (information retrieval); Computer science; Learning to rank; Information retrieval; Rank (graph theory); Preference; Artificial intelligence; Precision and recall; Macro; Key (lock); Latent variable; Machine learning; Natural language processing; Statistics; Mathematics","score_opus":0.03992445200966278,"score_gpt":0.24909795348301966,"score_spread":0.2091735014733569,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573576","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037426256,0.0032642218,0.9182228,0.0007609179,0.0005437395,0.00060580333,0.006553104,0.028682537,0.0039406936],"genre_scores_gemma":[0.3358349,0.0020012192,0.6132399,0.0005699758,0.0008195903,0.0006802679,0.026707478,0.0013709352,0.01877568],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981665,0.0005523836,0.00011949518,0.00040356978,0.00061080814,0.00014733608],"domain_scores_gemma":[0.99720275,0.00092533184,0.00032244535,0.0007407722,0.00065313594,0.00015570506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014601466,0.0021951275,0.0017819174,0.0040198006,0.0007274034,0.0015013174,0.0018190779,0.001005882,0.0045551835],"category_scores_gemma":[0.0070383064,0.0004897224,0.0011265095,0.0036525559,0.0005415972,0.003644584,0.001296489,0.0015111472,0.005189179],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055746944,0.0005706177,0.005152403,0.0008064361,0.00033753662,0.00028393738,0.00021889378,0.07365612,0.018409643,0.010322451,0.1196868,0.7699976],"study_design_scores_gemma":[0.00019817405,0.00030645917,0.0015487142,0.000036218924,0.00008032455,0.0004016752,0.0000955861,0.9485465,0.010934769,0.021719556,0.016035696,0.00009621074],"about_ca_topic_score_codex":0.0063598542,"about_ca_topic_score_gemma":0.01865264,"teacher_disagreement_score":0.0063598542,"about_ca_system_score_codex":0.00075602386,"about_ca_system_score_gemma":0.0017527455,"threshold_uncertainty_score":0.015238643},"labels":[],"label_agreement":null},{"id":"W4385573653","doi":"10.18653/v1/2022.emnlp-main.696","title":"Dictionary-Assisted Supervised Contrastive Learning","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; Charles Koch Foundation; Craig Newmark Philanthropies; John S. and James L. Knight Foundation; William and Flora Hewlett Foundation; Bill and Melinda Gates Foundation","keywords":"Computer science; Leverage (statistics); Natural language processing; Artificial intelligence; Entropy (arrow of time); Word (group theory); Cross entropy; Security token; Principle of maximum entropy; Linguistics","score_opus":0.02082510238209005,"score_gpt":0.2238134163181582,"score_spread":0.20298831393606817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573653","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044700522,0.00064155995,0.94981235,0.00027764766,0.00012155236,0.00012178094,0.00031811645,0.002291201,0.0017152177],"genre_scores_gemma":[0.606303,0.00030543553,0.3827296,0.00064820866,0.00029204975,0.00044416988,0.0031144593,0.00053875754,0.00562432],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9979469,0.0007748479,0.00012033908,0.0007274716,0.00031328545,0.00011708577],"domain_scores_gemma":[0.99349266,0.00396139,0.00045110798,0.0009986043,0.00092192966,0.00017431879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029930621,0.00141532,0.0013621113,0.0014392087,0.0007404901,0.0012018969,0.002195975,0.0014631372,0.0022180935],"category_scores_gemma":[0.010908813,0.00043005153,0.0010576396,0.0011314108,0.0014489589,0.0026454246,0.002748589,0.002827993,0.0013795123],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008653092,0.0008151134,0.0055387053,0.00043882694,0.00030923882,0.0002552262,0.0005999707,0.22132869,0.024707202,0.01417511,0.011679387,0.7192872],"study_design_scores_gemma":[0.000035785524,0.00013986863,0.00045007162,0.000015113291,0.000020990812,0.000059752336,0.000046979338,0.98158604,0.0059509343,0.010142633,0.0015349974,0.000016718202],"about_ca_topic_score_codex":0.0013195467,"about_ca_topic_score_gemma":0.003235851,"teacher_disagreement_score":0.0029930621,"about_ca_system_score_codex":0.00074801303,"about_ca_system_score_gemma":0.000980039,"threshold_uncertainty_score":0.015829027},"labels":[],"label_agreement":null},{"id":"W4385573710","doi":"10.18653/v1/2022.emnlp-main.803","title":"SPE: Symmetrical Prompt Enhancement for Fact Probing","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Task (project management); Computer science; Object (grammar); Artificial intelligence; Subject (documents); Symmetry (geometry); Task analysis; Natural language processing; Machine learning; Pattern recognition (psychology); Mathematics; Engineering","score_opus":0.04301873187770343,"score_gpt":0.2730614404365639,"score_spread":0.23004270855886044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573710","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14909478,0.0042537437,0.66688573,0.0022928633,0.0012141746,0.0012732045,0.021384122,0.14378132,0.009820128],"genre_scores_gemma":[0.4946448,0.00095687265,0.44281346,0.0017669753,0.00048249352,0.0012452028,0.04504959,0.0021740522,0.010866493],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983051,0.0006331385,0.00011927902,0.0006009352,0.0002286612,0.000112970745],"domain_scores_gemma":[0.99225426,0.00464824,0.00039822713,0.0017586212,0.00063413207,0.00030648746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028255414,0.0026176672,0.0008410841,0.0012949341,0.0005349322,0.001264608,0.0019698068,0.0017571679,0.008800063],"category_scores_gemma":[0.0148451915,0.0005114175,0.001253898,0.0010013317,0.00064874603,0.0052258093,0.0028117397,0.003865522,0.0061544618],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002003278,0.00085199816,0.009913418,0.0010625067,0.00020291432,0.00038225224,0.0012837293,0.0137259625,0.032553785,0.005351597,0.08500611,0.8476624],"study_design_scores_gemma":[0.000813387,0.0019417856,0.014145316,0.00027953033,0.00038609433,0.0010661524,0.0013242395,0.7511744,0.059921768,0.051197935,0.11748598,0.00026348437],"about_ca_topic_score_codex":0.0017212476,"about_ca_topic_score_gemma":0.0035607596,"teacher_disagreement_score":0.008800063,"about_ca_system_score_codex":0.00056513585,"about_ca_system_score_gemma":0.0015442432,"threshold_uncertainty_score":0.029439151},"labels":[],"label_agreement":null},{"id":"W4385573830","doi":"10.18653/v1/2022.findings-emnlp.76","title":"Faithful to the Document or to the World? Mitigating Hallucinations via Entity-Linked Knowledge in Abstractive Summarization","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Computer science; Knowledge base; Information retrieval; Source text; World Wide Web; Data science; Artificial intelligence","score_opus":0.023677054750677863,"score_gpt":0.2922954430037961,"score_spread":0.26861838825311823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573830","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11402344,0.0019921784,0.8735118,0.0028667755,0.000109425484,0.00019052265,0.0008978773,0.0029550716,0.0034529213],"genre_scores_gemma":[0.7831123,0.00110289,0.21135776,0.00044512938,0.00019288593,0.00010832351,0.0018624472,0.00021976707,0.0015985323],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9971041,0.0015539217,0.0001946492,0.00048546048,0.0005177119,0.00014411908],"domain_scores_gemma":[0.97422135,0.01687339,0.0026493988,0.0042832755,0.0016700261,0.00030255652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060828044,0.00088210695,0.00065434695,0.0022931013,0.0008309056,0.003186625,0.0013378042,0.0013286948,0.0017635605],"category_scores_gemma":[0.043280512,0.00043866696,0.0006002345,0.0018097466,0.0015806186,0.010429583,0.003097533,0.0026778562,0.0007414911],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013026588,0.00037994006,0.014763743,0.0015128353,0.00046480255,0.0012824581,0.0110554695,0.12014358,0.032453362,0.045657568,0.012276324,0.7587072],"study_design_scores_gemma":[0.0001151191,0.0005258038,0.008296876,0.0004241444,0.0006092795,0.0007749973,0.004573898,0.730994,0.048521943,0.17321612,0.031747807,0.00020005894],"about_ca_topic_score_codex":0.0026073435,"about_ca_topic_score_gemma":0.003874572,"teacher_disagreement_score":0.0060828044,"about_ca_system_score_codex":0.0007679207,"about_ca_system_score_gemma":0.0008248133,"threshold_uncertainty_score":0.032169342},"labels":[],"label_agreement":null},{"id":"W4385573851","doi":"10.18653/v1/2022.emnlp-main.521","title":"Text Style Transferring via Adversarial Masking and Styled Filling","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment","keywords":"Computer science; Adversarial system; Natural language processing; Masking (illustration); Consistency (knowledge bases); Artificial intelligence; Word (group theory); Entropy (arrow of time); Orthographic projection; Speech recognition; Linguistics","score_opus":0.014375260852957908,"score_gpt":0.21107195243731916,"score_spread":0.19669669158436126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573851","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.070773005,0.0005150906,0.9230458,0.00030685755,0.00020525508,0.0001716371,0.0002677266,0.0026531464,0.0020615074],"genre_scores_gemma":[0.77092725,0.000571166,0.21475102,0.0005264368,0.0003113887,0.00036329177,0.0010830545,0.00036407216,0.011102324],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99927956,0.00022071568,0.000043622567,0.00022828164,0.0001619607,0.00006582585],"domain_scores_gemma":[0.998214,0.0008618306,0.00019632923,0.0004975068,0.00015398997,0.00007641114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015346835,0.0013700527,0.0008252135,0.00069370033,0.00035821804,0.00069967617,0.0011600278,0.0010154741,0.0022156506],"category_scores_gemma":[0.0046382206,0.00035001323,0.0011452395,0.0006378988,0.0010165578,0.0015275867,0.0013694565,0.0014762185,0.0014614334],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062155304,0.00025693426,0.002943852,0.00026548922,0.00015217482,0.0005801824,0.00051367626,0.4061943,0.061934363,0.011540091,0.007933841,0.5070635],"study_design_scores_gemma":[0.00001863078,0.00012361926,0.000568568,0.00001332826,0.000022456386,0.00015191616,0.00002729591,0.9770123,0.013458603,0.0066532088,0.0019267363,0.00002336857],"about_ca_topic_score_codex":0.0007379076,"about_ca_topic_score_gemma":0.00081095856,"teacher_disagreement_score":0.0022156506,"about_ca_system_score_codex":0.00034840158,"about_ca_system_score_gemma":0.00047104413,"threshold_uncertainty_score":0.008116305},"labels":[],"label_agreement":null},{"id":"W4385573871","doi":"10.18653/v1/2022.emnlp-main.733","title":"Pneg: Prompt-based Negative Response Generation for Dialogue Response Selection Task","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Adversarial system; Computer science; Selection (genetic algorithm); Task (project management); Context (archaeology); Scalability; Artificial intelligence; Machine learning; Code (set theory); Language model; Model selection; Natural language processing; Programming language","score_opus":0.039145597082553156,"score_gpt":0.2701493714977883,"score_spread":0.23100377441523512,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573871","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029156897,0.00061396236,0.9174955,0.00055756955,0.0007474084,0.0010443265,0.0019329846,0.043561786,0.0048896507],"genre_scores_gemma":[0.37744412,0.0003158044,0.5864005,0.0013968203,0.0003968198,0.0027938925,0.006752415,0.0031199886,0.02137958],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9954129,0.0029888721,0.00012658155,0.00082950346,0.00045613342,0.00018595356],"domain_scores_gemma":[0.9935715,0.0041567893,0.000240767,0.0010173005,0.00074089353,0.0002727324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004985346,0.002696079,0.00110032,0.00068167184,0.0006153065,0.0010025415,0.002235696,0.0020452675,0.012067097],"category_scores_gemma":[0.014135678,0.00055922725,0.0008217926,0.00038938626,0.000954097,0.0018161744,0.002890572,0.0025496096,0.008868731],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004027722,0.0010656596,0.0034822323,0.0015184084,0.00024591907,0.0009219398,0.0020307547,0.06699344,0.11257535,0.011937071,0.118004285,0.6771973],"study_design_scores_gemma":[0.00028063112,0.0005736937,0.0010035465,0.00006882043,0.00006113618,0.00045414083,0.00033167892,0.8991274,0.047703702,0.016514646,0.03376137,0.00011927028],"about_ca_topic_score_codex":0.00091051834,"about_ca_topic_score_gemma":0.0015973374,"teacher_disagreement_score":0.012067097,"about_ca_system_score_codex":0.0005842313,"about_ca_system_score_gemma":0.0008652358,"threshold_uncertainty_score":0.040368497},"labels":[],"label_agreement":null},{"id":"W4385573948","doi":"10.18653/v1/2022.findings-emnlp.440","title":"WordTies: Measuring Word Associations in Language Models via Constrained Sampling","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Word (group theory); Natural language processing; Associative property; Artificial intelligence; Similarity (geometry); Proxy (statistics); Word Association; Language model; Property (philosophy); Cosine similarity; Machine learning; Linguistics; Mathematics; Pattern recognition (psychology)","score_opus":0.07116683802841292,"score_gpt":0.26913469582935595,"score_spread":0.19796785780094303,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573948","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06677966,0.0003595331,0.93015236,0.00017226322,0.000034915454,0.0001827545,0.00047608506,0.0011130229,0.0007293485],"genre_scores_gemma":[0.5819527,0.00036546896,0.4119402,0.00020161424,0.00012271109,0.00072123966,0.0032292074,0.0004479391,0.0010189054],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99495524,0.0025253466,0.00035418375,0.00116093,0.0008272799,0.00017698431],"domain_scores_gemma":[0.97295356,0.021443417,0.0017998051,0.0024114426,0.00095708075,0.0004345776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006600353,0.0014128557,0.0013647961,0.0033528542,0.0009968394,0.0027587728,0.0019577937,0.0018326098,0.00263996],"category_scores_gemma":[0.045313083,0.0007118505,0.0013348864,0.0031885158,0.0017718746,0.0063615777,0.003650138,0.0024955785,0.00085321476],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014579436,0.000640176,0.06616439,0.0010351357,0.0014685933,0.00044187796,0.0038759327,0.2326994,0.017717987,0.091187485,0.008233514,0.57507765],"study_design_scores_gemma":[0.00008487708,0.00020724081,0.0052971994,0.000061655184,0.00007772138,0.00020390897,0.00036382233,0.8606651,0.003842871,0.12678093,0.0023489564,0.00006580934],"about_ca_topic_score_codex":0.0033329828,"about_ca_topic_score_gemma":0.005141667,"teacher_disagreement_score":0.006600353,"about_ca_system_score_codex":0.0009286565,"about_ca_system_score_gemma":0.0013164217,"threshold_uncertainty_score":0.034906387},"labels":[],"label_agreement":null},{"id":"W4385573955","doi":"10.18653/v1/2022.emnlp-main.663","title":"Learning with Rejection for Abstractive Text Summarization","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"University of Manchester; Cancer Research UK","keywords":"Automatic summarization; Computer science; Artificial intelligence; Inference; Set (abstract data type); Natural language processing; Training set; Baseline (sea); Paraphrase; Machine learning; Decoding methods; Hallucinating","score_opus":0.014860129732340717,"score_gpt":0.23026875462342877,"score_spread":0.21540862489108806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573955","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027055949,0.00090596516,0.9658157,0.0004485054,0.000103441525,0.00012427682,0.0001773786,0.004069987,0.0012987232],"genre_scores_gemma":[0.61276305,0.00075873034,0.3737474,0.00069161627,0.0004520446,0.00044609202,0.0028962004,0.000875216,0.0073696286],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99736494,0.001398432,0.00017304094,0.0005285471,0.0004141633,0.000120980745],"domain_scores_gemma":[0.991916,0.005076246,0.00076065515,0.00092661375,0.0011189611,0.00020151002],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004600294,0.0014705537,0.0013258032,0.0011456383,0.000523895,0.0016523194,0.0020659724,0.0016032486,0.0021866134],"category_scores_gemma":[0.017135125,0.0004342201,0.0010282898,0.00077814603,0.0008403483,0.0028390964,0.0014289097,0.002770938,0.0015868631],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011766913,0.00035937867,0.0024844194,0.0006036305,0.00027210533,0.00022568817,0.0008676994,0.27842733,0.025199981,0.012154887,0.010901396,0.6673268],"study_design_scores_gemma":[0.000054559252,0.00024855093,0.00038728342,0.000029685398,0.00005387782,0.000065624816,0.00006866419,0.97792655,0.007873834,0.009739425,0.0035235882,0.000028347678],"about_ca_topic_score_codex":0.0017108717,"about_ca_topic_score_gemma":0.0024379082,"teacher_disagreement_score":0.004600294,"about_ca_system_score_codex":0.0007405458,"about_ca_system_score_gemma":0.00076152646,"threshold_uncertainty_score":0.024329007},"labels":[],"label_agreement":null},{"id":"W4385574022","doi":"10.18653/v1/2022.emnlp-main.694","title":"Human Guided Exploitation of Interpretable Attention Patterns in Summarization and Topic Segmentation","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of the Fraser Valley; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; University of British Columbia; Institute for Computing, Information and Cognitive Systems","keywords":"Automatic summarization; Computer science; Transformer; Segmentation; Artificial intelligence; Pipeline (software); Machine learning; Engineering; Voltage","score_opus":0.02972887344622373,"score_gpt":0.2780526593413482,"score_spread":0.24832378589512447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385574022","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060509034,0.00041519757,0.93095785,0.00047173447,0.000055247863,0.00013787927,0.00019351598,0.004790176,0.0024693452],"genre_scores_gemma":[0.7692878,0.00028338525,0.22552146,0.00024507602,0.000058064947,0.00013796153,0.0005217713,0.0004876327,0.0034569749],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99915504,0.00030153655,0.000033733875,0.00031976704,0.000105715815,0.0000842084],"domain_scores_gemma":[0.9964378,0.0023686276,0.00021871562,0.00049498456,0.0003531691,0.00012682927],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021310283,0.0013403326,0.0007598708,0.0010343465,0.00048523585,0.0016149676,0.0015517813,0.0011677436,0.0029808201],"category_scores_gemma":[0.009150731,0.0006286411,0.0009441907,0.00071540393,0.00096905185,0.0028944707,0.0015750627,0.001729421,0.0010865948],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011028609,0.00037257888,0.008523594,0.000622937,0.0002772674,0.00038165922,0.003507252,0.14469911,0.13486362,0.018149506,0.006253847,0.68124574],"study_design_scores_gemma":[0.000037223697,0.00017599511,0.0015988939,0.00002920724,0.00007948041,0.0001237494,0.00022390284,0.9495684,0.025730621,0.019511154,0.0028869556,0.000034460154],"about_ca_topic_score_codex":0.004406277,"about_ca_topic_score_gemma":0.008115192,"teacher_disagreement_score":0.004406277,"about_ca_system_score_codex":0.0008772752,"about_ca_system_score_gemma":0.0010586033,"threshold_uncertainty_score":0.011270106},"labels":[],"label_agreement":null},{"id":"W4385574133","doi":"10.18653/v1/2022.findings-emnlp.240","title":"AlphaTuning: Quantization-Aware Parameter-Efficient Adaptation of Large-Scale Pre-Trained Language Models","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Quantization (signal processing); Adaptation (eye); Language model; Scale (ratio); Artificial intelligence; Natural language processing; Speech recognition; Psychology; Algorithm; Physics; Quantum mechanics; Neuroscience","score_opus":0.028931766512733194,"score_gpt":0.2697405083380179,"score_spread":0.2408087418252847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385574133","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06477676,0.0030037917,0.9050478,0.00033534717,0.0007788535,0.00015459806,0.000669206,0.021201912,0.0040317723],"genre_scores_gemma":[0.6582957,0.0010557035,0.32828993,0.00039359622,0.00023074489,0.00034490973,0.0028239377,0.0024233386,0.006142207],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993319,0.0002125231,0.000049764167,0.00022567715,0.0001139464,0.000066145585],"domain_scores_gemma":[0.99771047,0.0011694906,0.000088156004,0.0005166655,0.0004134692,0.000101782694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014778599,0.0011654483,0.0010155516,0.0007727705,0.0004055452,0.0009658837,0.0019479943,0.0010648252,0.0035094332],"category_scores_gemma":[0.007215063,0.0006451022,0.0006353814,0.0008925586,0.0004917282,0.0019516151,0.0018317045,0.0021456124,0.0024037254],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000638967,0.00029411307,0.0016332722,0.0001740438,0.00016549969,0.00020033619,0.00030020182,0.13574955,0.037005518,0.0019601248,0.017817799,0.8040606],"study_design_scores_gemma":[0.00004593395,0.00006570848,0.00059507013,0.000015123216,0.00003465349,0.000055840414,0.00004717632,0.9844289,0.010090768,0.0024040223,0.0021931916,0.000023556811],"about_ca_topic_score_codex":0.007455328,"about_ca_topic_score_gemma":0.0111510465,"teacher_disagreement_score":0.007455328,"about_ca_system_score_codex":0.00060624356,"about_ca_system_score_gemma":0.0008919886,"threshold_uncertainty_score":0.014823854},"labels":[],"label_agreement":null},{"id":"W4385574234","doi":"10.18653/v1/2022.emnlp-industry.44","title":"Bringing the State-of-the-Art to Customers: A Neural Agent Assistant Framework for Customer Service Support","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Vector Institute; McGill University; PricewaterhouseCoopers (Canada); Queen's University","funders":"Vector Institute; University of Pennsylvania","keywords":"Customer service; Computer science; State (computer science); Service (business); Management; Artificial intelligence; Operations research; Engineering; Business; Marketing; Programming language; Economics","score_opus":0.030142063854225343,"score_gpt":0.27560313207799236,"score_spread":0.245461068223767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385574234","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03335379,0.001639652,0.9405028,0.0027190505,0.00028015365,0.000111071386,0.00023805167,0.004884953,0.016270496],"genre_scores_gemma":[0.71686035,0.0009791877,0.26139,0.0006351408,0.00017363684,0.00017279256,0.0003963387,0.00025436433,0.019138198],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996985,0.00011312585,0.000014805268,0.000068712645,0.00006005617,0.000044804117],"domain_scores_gemma":[0.9995198,0.00024568834,0.000030398218,0.000050092574,0.00010695541,0.00004706408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000935397,0.00043722393,0.00038094877,0.0004520193,0.00057910616,0.0013372666,0.0017261222,0.0012784554,0.004675778],"category_scores_gemma":[0.002358516,0.00030745481,0.00041319025,0.00036737425,0.0004601526,0.002405735,0.0016604268,0.001664448,0.00155668],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066132506,0.000668748,0.0021300185,0.0002757703,0.00019758765,0.0004022714,0.00086078676,0.22822581,0.007518172,0.07328992,0.0319362,0.6538334],"study_design_scores_gemma":[0.000012614719,0.000028139955,0.0001259489,0.000010846415,0.00002260477,0.000026566511,0.00004880303,0.97717893,0.0012871941,0.016412897,0.0048360666,0.00000946543],"about_ca_topic_score_codex":0.007575212,"about_ca_topic_score_gemma":0.014273539,"teacher_disagreement_score":0.007575212,"about_ca_system_score_codex":0.0005889898,"about_ca_system_score_gemma":0.0007203528,"threshold_uncertainty_score":0.015642047},"labels":[],"label_agreement":null},{"id":"W4385574290","doi":"10.18653/v1/2022.emnlp-main.597","title":"AfriCLIRMatrix: Enabling Cross-Lingual Information Retrieval for African Languages","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Computer science; Relevance (law); Information retrieval; Point (geometry); Natural language processing; Test (biology); Range (aeronautics); Feature (linguistics); Languages of Africa; Artificial intelligence; World Wide Web; Linguistics","score_opus":0.01964899349743088,"score_gpt":0.2931536038064133,"score_spread":0.27350461030898243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385574290","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09689583,0.009147762,0.081846416,0.0036870071,0.0013461126,0.002861297,0.49110267,0.25882706,0.054285865],"genre_scores_gemma":[0.10239277,0.0016747043,0.1900773,0.0011199523,0.00025045723,0.0022386338,0.6823101,0.0059090075,0.014027075],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99768984,0.0007453219,0.00020000746,0.00055436394,0.0005218857,0.00028855208],"domain_scores_gemma":[0.99750346,0.00076308264,0.00014556864,0.0007891596,0.00047058676,0.00032822046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033893161,0.0020765045,0.0013354104,0.006029203,0.0016460308,0.0030639134,0.0023048357,0.0011942731,0.016174521],"category_scores_gemma":[0.010244801,0.00065096724,0.0014565012,0.005823887,0.00060440804,0.0072651026,0.009019435,0.0018220268,0.016233357],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077673164,0.00046554944,0.0064536342,0.002016479,0.00025587494,0.0004469541,0.0016061638,0.0023192116,0.010784941,0.004219317,0.7638241,0.20683107],"study_design_scores_gemma":[0.0005633594,0.00047863295,0.018613687,0.0006212309,0.00020741604,0.0009110531,0.0032744377,0.07112093,0.025402296,0.011293794,0.86725676,0.000256481],"about_ca_topic_score_codex":0.017165523,"about_ca_topic_score_gemma":0.03250469,"teacher_disagreement_score":0.017165523,"about_ca_system_score_codex":0.0010497986,"about_ca_system_score_gemma":0.0020520762,"threshold_uncertainty_score":0.054109156},"labels":[],"label_agreement":null},{"id":"W4385574336","doi":"10.18653/v1/2022.findings-emnlp.363","title":"Improving Generalization of Pre-trained Language Models via Stochastic Weight Averaging","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Université de Montréal","funders":"","keywords":"Generalization; Computer science; Language model; Artificial intelligence; Flatness (cosmology); Computation; Distillation; Machine learning; Convergence (economics); Algorithm; Mathematics","score_opus":0.010534587294332803,"score_gpt":0.22110124622583868,"score_spread":0.2105666589315059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385574336","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040870313,0.0005892828,0.9540103,0.0002823457,0.00006333121,0.000049627623,0.00013820663,0.002834052,0.0011625697],"genre_scores_gemma":[0.6984598,0.00081693335,0.29227254,0.0006441017,0.00018684193,0.00033066623,0.001631746,0.0009233901,0.0047339383],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991043,0.00029508251,0.00007342481,0.00028073342,0.00016409661,0.000082380866],"domain_scores_gemma":[0.99684227,0.0019530505,0.00017263913,0.00061986805,0.0003263573,0.000085802945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002236412,0.0021012062,0.00161387,0.0009844013,0.00055744185,0.0011401387,0.0018672479,0.0014622346,0.0018744282],"category_scores_gemma":[0.009707629,0.00093185564,0.0014735662,0.00095760584,0.00088593614,0.0035389136,0.0022640086,0.0031811497,0.0012033478],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008597302,0.00011115809,0.0008141215,0.00011943289,0.00015848367,0.00010200362,0.00017035197,0.7970065,0.00938088,0.0072820433,0.0027296855,0.18203938],"study_design_scores_gemma":[0.0000054452507,0.000020870337,0.00008632335,0.0000050217805,0.000010071543,0.000012134938,0.0000075352646,0.9948226,0.0009268888,0.0038378392,0.00025921775,0.0000060267566],"about_ca_topic_score_codex":0.007859498,"about_ca_topic_score_gemma":0.012624717,"teacher_disagreement_score":0.007859498,"about_ca_system_score_codex":0.0009224016,"about_ca_system_score_gemma":0.0013390376,"threshold_uncertainty_score":0.015627503},"labels":[],"label_agreement":null},{"id":"W4385574438","doi":"10.18653/v1/2022.finnlp-1.8","title":"Learning Better Intent Representations for Financial Open Intent Classification","year":2022,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Benchmark (surveying); Artificial intelligence; Machine learning; Domain (mathematical analysis); Open domain; Decision boundary; Fine-tuning; Natural language processing; Question answering; Support vector machine","score_opus":0.10664158920130207,"score_gpt":0.3312236231784866,"score_spread":0.22458203397718451,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385574438","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23456733,0.0026131826,0.73183036,0.0016914209,0.0005025358,0.00036065702,0.0027207094,0.018844623,0.0068691676],"genre_scores_gemma":[0.7662303,0.00051384134,0.2159247,0.0005796241,0.00027762546,0.00027441693,0.011119129,0.00046869027,0.00461169],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99903584,0.00030401108,0.000079990794,0.00024077184,0.0001740567,0.00016523547],"domain_scores_gemma":[0.9978605,0.0010916182,0.00016365008,0.00042925827,0.00033123614,0.00012362067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019405901,0.0017621916,0.0008153777,0.0023760712,0.0005470711,0.0016957807,0.0012797194,0.0015005476,0.0027345198],"category_scores_gemma":[0.0060686166,0.00036151957,0.0014880097,0.0012357556,0.0004903614,0.004227984,0.0017477666,0.0031409827,0.002553203],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055381865,0.00097556895,0.016496304,0.00040011696,0.0001326916,0.00032455777,0.0008975953,0.049014673,0.016975546,0.008846555,0.040131606,0.86525106],"study_design_scores_gemma":[0.000051761064,0.00018125563,0.0026423607,0.00008807406,0.00007061464,0.00016967932,0.0004123459,0.9546612,0.0077287843,0.02707593,0.0068769725,0.00004095051],"about_ca_topic_score_codex":0.002271709,"about_ca_topic_score_gemma":0.0034245881,"teacher_disagreement_score":0.0027345198,"about_ca_system_score_codex":0.000630701,"about_ca_system_score_gemma":0.000825302,"threshold_uncertainty_score":0.010262966},"labels":[],"label_agreement":null},{"id":"W4385595332","doi":"10.1145/3573128.3609347","title":"AI-powered Resume-Job matching","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Leverage (statistics); Artificial intelligence; Matching (statistics); Sentence; Natural language processing; Similarity (geometry); Transformer; Discriminative model; Representation (politics); Information retrieval; Machine learning; Image (mathematics)","score_opus":0.031796046000300754,"score_gpt":0.28167254926802715,"score_spread":0.2498765032677264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385595332","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42967224,0.0031894678,0.49321967,0.0016571565,0.00060585554,0.0007275884,0.003431973,0.015927907,0.051568232],"genre_scores_gemma":[0.8948536,0.00051621377,0.08379956,0.00034554116,0.00012398705,0.0001242382,0.003948597,0.00018760037,0.016100658],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925035,0.00015100634,0.000045310117,0.00020051525,0.00023963199,0.000113180104],"domain_scores_gemma":[0.9984516,0.00044891104,0.00017048168,0.00042860562,0.00037273325,0.00012767708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013967601,0.0005795675,0.00051511073,0.0012989683,0.0004978523,0.0012745151,0.0015224549,0.0008737918,0.0062192176],"category_scores_gemma":[0.0056295097,0.00020102832,0.0004391276,0.0014533554,0.0004350158,0.003001702,0.0011658259,0.000966624,0.0032326533],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076901173,0.00079999387,0.020679945,0.00048244908,0.0001290343,0.00037168525,0.0003016778,0.1449167,0.017798869,0.01581105,0.024329677,0.77361],"study_design_scores_gemma":[0.000045330107,0.00024924215,0.007582291,0.000029105873,0.00003921871,0.00022933904,0.00022637186,0.9508149,0.015269489,0.012653058,0.012829385,0.000032288688],"about_ca_topic_score_codex":0.008883581,"about_ca_topic_score_gemma":0.01167139,"teacher_disagreement_score":0.008883581,"about_ca_system_score_codex":0.0010646979,"about_ca_system_score_gemma":0.0018180402,"threshold_uncertainty_score":0.0208053},"labels":[],"label_agreement":null},{"id":"W4385612789","doi":"10.1145/3539618.3591925","title":"SIGIR 2023 Workshop on Retrieval Enhanced Machine Learning (REML @ SIGIR 2023)","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada)","funders":"","keywords":"Computer science; Machine learning; Artificial intelligence; ENCODE; Question answering; Context (archaeology); Robustness (evolution); Information retrieval","score_opus":0.03283412424060274,"score_gpt":0.27655578519061547,"score_spread":0.24372166095001274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385612789","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021388687,0.20110825,0.5052229,0.10085721,0.06948334,0.0025270213,0.010708201,0.024728656,0.06397573],"genre_scores_gemma":[0.07852136,0.05658193,0.5441977,0.034250766,0.033442352,0.003182943,0.039960533,0.005262069,0.20460032],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.986845,0.006354034,0.0006243329,0.0019180203,0.0032159628,0.0010426036],"domain_scores_gemma":[0.97792315,0.011179286,0.0005581294,0.0028042435,0.004708117,0.0028270124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0376346,0.0039240387,0.0044920784,0.0053698854,0.0015603066,0.0074743526,0.0076001887,0.008031345,0.023709318],"category_scores_gemma":[0.026730072,0.0013775587,0.0023323048,0.0037430339,0.002601418,0.013080022,0.0060740006,0.009218978,0.02342568],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038672172,0.0004274769,0.00032567157,0.0005307544,0.0001667602,0.00014685608,0.0002420557,0.0022506618,0.002460695,0.006646497,0.7612077,0.22520816],"study_design_scores_gemma":[0.00034737232,0.0006857322,0.0020927335,0.00060658046,0.00020995362,0.00069956615,0.000506514,0.051516127,0.0066580265,0.039781906,0.8966783,0.00021705354],"about_ca_topic_score_codex":0.01257276,"about_ca_topic_score_gemma":0.015636712,"teacher_disagreement_score":0.0376346,"about_ca_system_score_codex":0.0038657133,"about_ca_system_score_gemma":0.004626456,"threshold_uncertainty_score":0.19903314},"labels":[],"label_agreement":null},{"id":"W4385685557","doi":"10.2139/ssrn.4533669","title":"Multimodal Dialogue Modeling: Simultaneous Intent Recognition and Response Generation","year":2023,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Speech recognition; Natural language processing; Artificial intelligence","score_opus":0.07046017678332024,"score_gpt":0.27458905974790626,"score_spread":0.20412888296458603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385685557","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024195813,0.00022487607,0.96810293,0.00032013073,0.00010669929,0.00015758835,0.00030564395,0.0043877214,0.0021985911],"genre_scores_gemma":[0.6274189,0.00024946345,0.3625819,0.00023324533,0.00024128957,0.00044391627,0.0013530969,0.0008003686,0.006677743],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998156,0.0009039262,0.00006732946,0.00045866368,0.000252489,0.00016152885],"domain_scores_gemma":[0.9972728,0.0018372798,0.00012890561,0.00030877555,0.00032164762,0.00013060671],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019240471,0.0012102633,0.0012225693,0.0008018547,0.0005037289,0.0017728278,0.0013915933,0.0014992716,0.0056977104],"category_scores_gemma":[0.007265161,0.0006630016,0.0011855118,0.0006015822,0.00037538473,0.0017470405,0.0016967707,0.0016649227,0.003457839],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018775017,0.0007708666,0.0032065606,0.00056499115,0.00024997152,0.0003855309,0.0015183295,0.10565851,0.081900194,0.010685075,0.014991775,0.77819073],"study_design_scores_gemma":[0.000027388272,0.000105199055,0.0005267347,0.000015742318,0.000039669932,0.000075966374,0.00010848648,0.98087746,0.010373897,0.006447789,0.0013764744,0.00002519214],"about_ca_topic_score_codex":0.0023663244,"about_ca_topic_score_gemma":0.0023942182,"teacher_disagreement_score":0.0056977104,"about_ca_system_score_codex":0.0003799316,"about_ca_system_score_gemma":0.000740902,"threshold_uncertainty_score":0.019060671},"labels":[],"label_agreement":null},{"id":"W4385688511","doi":"10.1145/3578337.3605136","title":"Perspectives on Large Language Models for Relevance Judgment","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":152,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Science Foundation","keywords":"Relevance (law); Perspective (graphical); Categorization; Computer science; Point (geometry); Compromise; Cognitive psychology; Psychology; Artificial intelligence; Sociology; Political science; Social science","score_opus":0.03503074496303086,"score_gpt":0.2930385171724492,"score_spread":0.25800777220941834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385688511","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015479181,0.009073187,0.88668036,0.060471475,0.0003983487,0.0001458004,0.0003051957,0.0013797939,0.026066693],"genre_scores_gemma":[0.6278574,0.0033070296,0.35493279,0.004682542,0.0029275254,0.0008475192,0.00049829355,0.0006766091,0.004270291],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.94592535,0.042749092,0.0015033085,0.0038909556,0.0050731106,0.00085809693],"domain_scores_gemma":[0.6764201,0.2811487,0.00691776,0.024959132,0.007913221,0.0026411219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07345315,0.001647202,0.002302094,0.0046816845,0.0027299542,0.014315395,0.0049077873,0.0057553137,0.0063704033],"category_scores_gemma":[0.16614789,0.0015305054,0.0018020578,0.0029826474,0.012353361,0.029320093,0.0072641806,0.009937408,0.0021383038],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001170609,0.00007078386,0.0009936986,0.00031618442,0.00011130567,0.00008174834,0.0025947243,0.009490547,0.0008249442,0.9440781,0.0039639,0.03735708],"study_design_scores_gemma":[0.000031351687,0.000027390735,0.00035846487,0.000081098595,0.000022872486,0.000056826473,0.00027198257,0.048352472,0.00046027006,0.94326943,0.0070259515,0.000041864423],"about_ca_topic_score_codex":0.0039215568,"about_ca_topic_score_gemma":0.0029704662,"teacher_disagreement_score":0.07345315,"about_ca_system_score_codex":0.0049744276,"about_ca_system_score_gemma":0.003059765,"threshold_uncertainty_score":0.388462},"labels":[],"label_agreement":null},{"id":"W4385701337","doi":"10.1007/978-3-031-39831-5_22","title":"Exploring Dialog Act Recognition in Open Domain Conversational Agents","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Dialog box; Computer science; Classifier (UML); Open domain; Dialog system; Artificial intelligence; Support vector machine; Domain (mathematical analysis); Baseline (sea); Natural language processing; Machine learning; Speech recognition; World Wide Web; Question answering","score_opus":0.1822798249837222,"score_gpt":0.2925068874468192,"score_spread":0.11022706246309699,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385701337","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15483809,0.0013857639,0.8321066,0.0005082978,0.00010865731,0.00013507785,0.0005036707,0.003635698,0.0067781513],"genre_scores_gemma":[0.78517956,0.00029581253,0.20776697,0.00011948124,0.000064234446,0.00010714531,0.0011617488,0.00028065816,0.0050244075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99907804,0.0004156596,0.0000361891,0.00025179697,0.000112570204,0.000105789055],"domain_scores_gemma":[0.9977016,0.0018768212,0.00008045079,0.00012627363,0.00012971023,0.000085127554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011500848,0.0007458907,0.0009039686,0.000580387,0.0007256005,0.0026821245,0.0012460201,0.00125687,0.0037378077],"category_scores_gemma":[0.004148556,0.00051531085,0.00087832025,0.0005098347,0.00050528604,0.002728734,0.001978709,0.0018262768,0.0011264973],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016540776,0.0007759254,0.0079343375,0.0005856426,0.00030130445,0.0007071377,0.0046655196,0.096789256,0.060414866,0.030486647,0.0095273545,0.78615797],"study_design_scores_gemma":[0.000016967244,0.00007936852,0.0015323893,0.00003082408,0.00004141936,0.00011132693,0.00082804833,0.96004343,0.00902204,0.024756525,0.0035134393,0.000024203584],"about_ca_topic_score_codex":0.0034844277,"about_ca_topic_score_gemma":0.0037049023,"teacher_disagreement_score":0.0037378077,"about_ca_system_score_codex":0.0005084242,"about_ca_system_score_gemma":0.00053363765,"threshold_uncertainty_score":0.01250422},"labels":[],"label_agreement":null},{"id":"W4385718129","doi":"10.18653/v1/2023.wassa-1.37","title":"Exploration of Contrastive Learning Strategies toward more Robust Stance Detection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Artificial intelligence; Adversarial system; Sentence; Embedding; Transformer; Natural language processing; Task (project management); Machine learning","score_opus":0.07783445247058172,"score_gpt":0.27766335827664207,"score_spread":0.19982890580606033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385718129","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15099218,0.00078642904,0.84197474,0.00071287324,0.000091499256,0.00013731133,0.00016867595,0.002383075,0.0027532035],"genre_scores_gemma":[0.86261517,0.0002048102,0.13384683,0.00043493122,0.0000957768,0.00010873032,0.00040920437,0.00021461264,0.0020699957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921143,0.0002968685,0.000041079413,0.00027634716,0.000112313275,0.000061939805],"domain_scores_gemma":[0.9974899,0.0015472301,0.00022139394,0.0003362331,0.00027364746,0.00013159859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021923387,0.0012897969,0.0007891061,0.0009610948,0.0003691331,0.001011312,0.0012823793,0.0010466086,0.0018124535],"category_scores_gemma":[0.0063411267,0.00031154236,0.00069276715,0.00042227635,0.00097328087,0.0024508666,0.001659974,0.0019409262,0.0010263235],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009487242,0.0010066913,0.008712229,0.00032638063,0.0002814844,0.00058159285,0.0006903162,0.2820617,0.12022941,0.019750109,0.006319474,0.559092],"study_design_scores_gemma":[0.000026797843,0.00018177507,0.00037146144,0.000013095334,0.000021154468,0.00007355369,0.000050989293,0.9795975,0.008626448,0.010332282,0.00069429673,0.000010665252],"about_ca_topic_score_codex":0.0009159728,"about_ca_topic_score_gemma":0.0016281233,"teacher_disagreement_score":0.0021923387,"about_ca_system_score_codex":0.0006161474,"about_ca_system_score_gemma":0.00065594417,"threshold_uncertainty_score":0.0115942955},"labels":[],"label_agreement":null},{"id":"W4385718159","doi":"10.18653/v1/2023.sigmorphon-1.23","title":"An Ensembled Encoder-Decoder System for Interlinear Glossed Text","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Security token; Task (project management); Sequence (biology); Encoder; Resource (disambiguation); Set (abstract data type); Test set; Artificial intelligence; Speech recognition; Natural language processing; Machine learning; Computer security; Engineering; Programming language","score_opus":0.038773114858213505,"score_gpt":0.29329555004059615,"score_spread":0.25452243518238266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385718159","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035769556,0.0030225895,0.75601983,0.0014911805,0.0019386368,0.0005088338,0.017428227,0.1688061,0.015015095],"genre_scores_gemma":[0.2537645,0.0011138124,0.6429636,0.0011322127,0.0006656493,0.00048046082,0.06320213,0.0054382863,0.031239424],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99897814,0.00020368131,0.00008821139,0.000426959,0.0002005637,0.00010242111],"domain_scores_gemma":[0.9980514,0.0006351512,0.00007578673,0.00060284976,0.00051956286,0.000115291165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013577389,0.001990993,0.0012585891,0.0013504506,0.0007878067,0.0016299302,0.0019455373,0.0013859029,0.015731143],"category_scores_gemma":[0.005350123,0.0007547098,0.0011522993,0.0011518521,0.0003980719,0.0047987797,0.0028691606,0.0028725648,0.017507564],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043549435,0.00022060124,0.0015196509,0.00037759126,0.00023201019,0.00043224715,0.0002834846,0.019443393,0.020559426,0.008428214,0.12098492,0.8270829],"study_design_scores_gemma":[0.0001129938,0.00018545716,0.0010743424,0.00007489318,0.00017592455,0.00050935807,0.00020002603,0.90602624,0.029333096,0.020089675,0.04212529,0.00009272438],"about_ca_topic_score_codex":0.008684877,"about_ca_topic_score_gemma":0.01979791,"teacher_disagreement_score":0.015731143,"about_ca_system_score_codex":0.00093529176,"about_ca_system_score_gemma":0.0015598432,"threshold_uncertainty_score":0.052625954},"labels":[],"label_agreement":null},{"id":"W4385732413","doi":"10.1109/tkde.2023.3303916","title":"XMQAs: Constructing Complex-Modified Question-Answering Dataset for Robust Question Understanding","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Knowledge and Data Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute","funders":"Science and Technology Commission of Shanghai Municipality; National Natural Science Foundation of China","keywords":"Computer science; Question answering; Robustness (evolution); Construct (python library); Semantics (computer science); Simple (philosophy); Artificial intelligence; Machine learning; Information retrieval; Natural language processing; Programming language","score_opus":0.15051792830781893,"score_gpt":0.3267831472600723,"score_spread":0.17626521895225336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385732413","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16839415,0.004817163,0.4124681,0.0034608075,0.0009519354,0.0064812223,0.32271573,0.07076926,0.009941677],"genre_scores_gemma":[0.11417629,0.0005218745,0.36520702,0.0011096037,0.00014133955,0.0036617727,0.51190436,0.000412882,0.0028648076],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9949043,0.0016728308,0.0007685109,0.0015072154,0.00091435417,0.00023282667],"domain_scores_gemma":[0.99225616,0.0028725038,0.00053309667,0.002080187,0.0018074667,0.00045058195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047216644,0.002148067,0.0010971627,0.005414131,0.001327712,0.0019452437,0.0041683675,0.0028754768,0.0042943773],"category_scores_gemma":[0.018872237,0.0004930668,0.0023639898,0.0034242298,0.0008997871,0.0044537806,0.004437661,0.003073293,0.0031984139],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013895408,0.002624623,0.03651801,0.006241007,0.0007466947,0.00105363,0.00402556,0.028665284,0.043389656,0.021445785,0.36846942,0.48543084],"study_design_scores_gemma":[0.0007635216,0.0011457554,0.04400479,0.0005688471,0.00040894642,0.001312518,0.0037341802,0.4841541,0.046732467,0.040407125,0.37637588,0.000391793],"about_ca_topic_score_codex":0.012933268,"about_ca_topic_score_gemma":0.016172487,"teacher_disagreement_score":0.012933268,"about_ca_system_score_codex":0.0016560357,"about_ca_system_score_gemma":0.0024184047,"threshold_uncertainty_score":0.025715947},"labels":[],"label_agreement":null},{"id":"W4385734218","doi":"10.18653/v1/2023.matching-1.7","title":"Knowledge-Augmented Language Model Prompting for Zero-Shot Knowledge Graph Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Question answering; Computer science; Shot (pellet); Knowledge graph; Graph; Task (project management); Zero (linguistics); Domain knowledge; Natural language processing; Artificial intelligence; Information retrieval; Theoretical computer science; Linguistics","score_opus":0.06291405592591279,"score_gpt":0.33947156824000024,"score_spread":0.27655751231408743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385734218","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035610616,0.0009284869,0.9017234,0.00069534493,0.00023085103,0.00033515034,0.0012759977,0.056210924,0.0029892998],"genre_scores_gemma":[0.5705777,0.00037146272,0.41664934,0.001094533,0.00015596263,0.00044594146,0.00458746,0.0010431887,0.0050744372],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979603,0.00090261636,0.00008417767,0.0006427499,0.00028038496,0.0001297107],"domain_scores_gemma":[0.9950204,0.00332325,0.00016404639,0.000845843,0.00043507825,0.00021143393],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020189236,0.001810062,0.0010759081,0.00080700935,0.0005530274,0.0012894283,0.0025004132,0.0024042479,0.008675625],"category_scores_gemma":[0.012030231,0.00052653626,0.000938855,0.00053983275,0.0009470495,0.0050144815,0.003551623,0.0033420683,0.0038705377],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017536682,0.000991486,0.0026860859,0.0018109928,0.00017014788,0.0008366502,0.0025888039,0.07591418,0.06530229,0.016760431,0.039668184,0.79151714],"study_design_scores_gemma":[0.00016515625,0.00048589546,0.00087644457,0.000069742804,0.00009284185,0.00033621097,0.00054002396,0.89249057,0.027997218,0.056365497,0.02049326,0.00008721094],"about_ca_topic_score_codex":0.0036598535,"about_ca_topic_score_gemma":0.0066327695,"teacher_disagreement_score":0.008675625,"about_ca_system_score_codex":0.00091769727,"about_ca_system_score_gemma":0.001375325,"threshold_uncertainty_score":0.029022872},"labels":[],"label_agreement":null},{"id":"W4385734257","doi":"10.18653/v1/2023.trustnlp-1.8","title":"Reliability Check: An Analysis of GPT-3’s Response to Sensitive Topics and Prompt Wording","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Mainstream; Consistency (knowledge bases); Reliability (semiconductor); Computer science; Reinforcement learning; Risk analysis (engineering); Work (physics); Simple (philosophy); Data science; Artificial intelligence; Engineering; Law; Political science; Epistemology; Power (physics); Medicine","score_opus":0.028798909874304987,"score_gpt":0.29581460532775733,"score_spread":0.26701569545345233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385734257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8723729,0.0004049,0.08595112,0.0025955115,0.0006553568,0.00078425166,0.0027773688,0.026514819,0.0079437755],"genre_scores_gemma":[0.9620785,0.000113683695,0.025761891,0.0012572751,0.00008124274,0.0006491797,0.0039581675,0.0033640559,0.002736059],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95508206,0.026728462,0.0026633292,0.0050924663,0.009298099,0.0011356679],"domain_scores_gemma":[0.5802977,0.34300864,0.012116804,0.033419102,0.028359178,0.0027985405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030952014,0.0018288264,0.001156873,0.0013879159,0.00089243706,0.0021858457,0.0023542743,0.0024846664,0.0032914055],"category_scores_gemma":[0.3180653,0.0007936315,0.0007237776,0.0011874427,0.0018431672,0.003062447,0.0034736437,0.005682676,0.0029445493],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01273195,0.002641883,0.3031015,0.00393148,0.0010770138,0.0061396887,0.089019395,0.09321018,0.07880874,0.009234651,0.10031999,0.29978356],"study_design_scores_gemma":[0.0006356098,0.0029290335,0.09599884,0.00070520095,0.0004014598,0.0019206331,0.011790251,0.76420224,0.0589158,0.012819365,0.048936486,0.0007451114],"about_ca_topic_score_codex":0.005679793,"about_ca_topic_score_gemma":0.003057623,"teacher_disagreement_score":0.030952014,"about_ca_system_score_codex":0.0015560492,"about_ca_system_score_gemma":0.0018747913,"threshold_uncertainty_score":0.16369188},"labels":[],"label_agreement":null},{"id":"W4385759351","doi":"10.3758/s13421-023-01449-9","title":"Instance theory predicts categorization decisions in the absence of categorical structure: A computational analysis of artificial grammar learning without a grammar","year":2023,"lang":"en","type":"article","venue":"Memory & Cognition","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Booth University College","funders":"","keywords":"Categorization; Psychology; Cognitive psychology; Grammar; Categorical variable; Cognition; Stimulus (psychology); Implicit learning; Phenomenon; Cognitive science; Linguistics; Artificial intelligence; Computer science; Machine learning; Epistemology","score_opus":0.031050715929445145,"score_gpt":0.2772898030497176,"score_spread":0.24623908712027248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385759351","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88024145,0.000089972964,0.11242239,0.0020814477,0.00004996186,0.00003055815,0.0001232739,0.00019214378,0.0047688135],"genre_scores_gemma":[0.98709095,0.00003949585,0.012101963,0.000108257795,0.000031093383,0.000021734373,0.00010949537,0.000039386166,0.0004576745],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993098,0.0003192711,0.000028264585,0.00016585395,0.00009675984,0.000080140635],"domain_scores_gemma":[0.977218,0.019913832,0.00062518945,0.0012563656,0.0004975706,0.00048906036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031323438,0.00030316872,0.001019651,0.00086069974,0.0006675247,0.0028819893,0.002085476,0.0017392931,0.0037739482],"category_scores_gemma":[0.021512702,0.000564668,0.0013030266,0.00058156706,0.0019122143,0.006018576,0.0013238683,0.0030043656,0.0003188532],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006464254,0.00055801,0.031441793,0.00022135722,0.00031253457,0.00041779727,0.0011701948,0.18056527,0.0061997036,0.72260857,0.0051982007,0.050660145],"study_design_scores_gemma":[0.00003681116,0.000030871866,0.002028485,0.00000650471,0.000025227706,0.00004836526,0.000069648624,0.6063229,0.0004279642,0.39084682,0.00014340733,0.000012987944],"about_ca_topic_score_codex":0.002192106,"about_ca_topic_score_gemma":0.0018847388,"teacher_disagreement_score":0.0037739482,"about_ca_system_score_codex":0.0009951139,"about_ca_system_score_gemma":0.0007903591,"threshold_uncertainty_score":0.01656562},"labels":[],"label_agreement":null},{"id":"W4385767747","doi":"10.24963/ijcai.2023/460","title":"On Conditional and Compositional Language Model Differentiable Prompting","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Modular design; Principle of compositionality; Generalization; Automatic summarization; Task (project management); Artificial intelligence; Word (group theory); Natural language processing; Embedding; Language model; Artificial neural network; Human–computer interaction; Programming language","score_opus":0.021424134020040075,"score_gpt":0.2602325658297621,"score_spread":0.23880843180972205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385767747","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05564792,0.00034626684,0.9373673,0.0003481116,0.00007374837,0.00005939133,0.00018987832,0.0034857404,0.0024816555],"genre_scores_gemma":[0.82046413,0.0003403467,0.1698323,0.00026816243,0.0000778954,0.00021244895,0.00069248606,0.00046319468,0.007649232],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994271,0.00026443953,0.000024003497,0.00017848019,0.000058769747,0.00004730522],"domain_scores_gemma":[0.99798524,0.0012505366,0.00011317639,0.0003910149,0.00018832689,0.0000717034],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001476727,0.00080582825,0.0006413082,0.00029700692,0.0002828562,0.0006656728,0.0013484855,0.0008774068,0.0037392837],"category_scores_gemma":[0.0064934413,0.00044585927,0.00065417925,0.0004275603,0.0010420191,0.0029943986,0.0015382356,0.0022170467,0.0010947034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028915133,0.00014413787,0.0010647088,0.00024258092,0.00005959329,0.00021693777,0.00048143746,0.724988,0.0138222175,0.040774845,0.0034559919,0.21446043],"study_design_scores_gemma":[0.000012253236,0.000049839982,0.00013123237,0.000006068852,0.000006588475,0.000024644032,0.000017535922,0.9741823,0.0018906476,0.022987274,0.0006829452,0.000008688547],"about_ca_topic_score_codex":0.0019958748,"about_ca_topic_score_gemma":0.0026110215,"teacher_disagreement_score":0.0037392837,"about_ca_system_score_codex":0.0006679253,"about_ca_system_score_gemma":0.0007603387,"threshold_uncertainty_score":0.012509167},"labels":[],"label_agreement":null},{"id":"W4385767844","doi":"10.24963/ijcai.2023/561","title":"KEST: Kernel Distance Based Efficient Self-Training for Improving Controllable Text Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Text generation; Generator (circuit theory); Bottleneck; Natural language generation; Fluency; Kernel (algebra); Artificial intelligence; Exploit; Language model; Process (computing); Machine learning; Natural language; Power (physics); Mathematics","score_opus":0.03806845852147484,"score_gpt":0.25414839389871313,"score_spread":0.2160799353772383,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385767844","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029353095,0.00031946093,0.96336526,0.00011513916,0.00006484662,0.0000681528,0.00009839278,0.005494122,0.0011215819],"genre_scores_gemma":[0.5999564,0.00025083552,0.3885623,0.0004038967,0.0000909286,0.0003254813,0.0015228149,0.0014616136,0.0074257767],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991966,0.000283679,0.000054118886,0.00021875674,0.00017278647,0.000074102434],"domain_scores_gemma":[0.99722964,0.0016461145,0.00018867181,0.00043291572,0.0003867593,0.0001159842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015543256,0.0010669845,0.0008500152,0.00076687423,0.00043582247,0.00065871375,0.0018063792,0.0011440265,0.0028981476],"category_scores_gemma":[0.006381124,0.00045366745,0.0007456285,0.0006138798,0.0008101769,0.0022105165,0.001917104,0.0017645295,0.0016930818],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038124112,0.00034797617,0.0019895073,0.0002396219,0.000093548144,0.00018309304,0.00040143586,0.41049534,0.025342848,0.0092899185,0.0077048494,0.54353064],"study_design_scores_gemma":[0.0000133848425,0.00004665438,0.000107821,0.0000052190194,0.0000053593035,0.000027746948,0.000013055034,0.9928941,0.0037950946,0.0024694188,0.00061527337,0.0000069046478],"about_ca_topic_score_codex":0.0019241633,"about_ca_topic_score_gemma":0.0033937567,"teacher_disagreement_score":0.0028981476,"about_ca_system_score_codex":0.0005901761,"about_ca_system_score_gemma":0.00083681283,"threshold_uncertainty_score":0.0096952915},"labels":[],"label_agreement":null},{"id":"W4385775400","doi":"","title":"Complex question answering: homogeneous or heterogeneous, which ensemble is better?","year":2014,"lang":"en","type":"article","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Homogeneous; Computer science; Question answering; Artificial intelligence; Natural language processing; Statistical physics; Physics","score_opus":0.0212606667660785,"score_gpt":0.23774858163454227,"score_spread":0.21648791486846378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385775400","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21974498,0.015333866,0.7114848,0.021041334,0.00096989225,0.0005807486,0.0048548896,0.003796612,0.02219281],"genre_scores_gemma":[0.80226356,0.004156959,0.16957602,0.0036041895,0.004289974,0.00031202025,0.009494769,0.0007953618,0.005507199],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99030405,0.003681215,0.00054586004,0.0035299752,0.0014614821,0.00047738108],"domain_scores_gemma":[0.9575369,0.027370587,0.0017406004,0.0060779294,0.004784938,0.00248905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013253337,0.0016973523,0.0034104802,0.0023876147,0.0012484328,0.007295046,0.0022743715,0.0029203522,0.012074272],"category_scores_gemma":[0.041888457,0.00055942213,0.0020338828,0.0022520367,0.0015924635,0.012260893,0.0038214147,0.0034009719,0.0037368827],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031033317,0.0011274924,0.037623852,0.002308764,0.0024715855,0.00030905724,0.0021074044,0.037385322,0.02284802,0.027074622,0.050508287,0.81313235],"study_design_scores_gemma":[0.00033197974,0.0009617326,0.034698058,0.0005544019,0.0019057166,0.0010726626,0.0024296811,0.53386074,0.015496854,0.35175395,0.056722883,0.00021139963],"about_ca_topic_score_codex":0.0016637533,"about_ca_topic_score_gemma":0.0015524685,"teacher_disagreement_score":0.013253337,"about_ca_system_score_codex":0.0010786683,"about_ca_system_score_gemma":0.0010608177,"threshold_uncertainty_score":0.07009113},"labels":[],"label_agreement":null},{"id":"W4385780698","doi":"10.1145/3613447","title":"Toward Best Practices for Training Multilingual Dense Retrieval Models","year":2023,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund; Compute Canada","keywords":"Computer science; Transformer; Language model; Architecture; Encoder; Relevance (law); Natural language processing; Artificial intelligence; Transfer of learning; Training set; Information retrieval; Variety (cybernetics); Data science","score_opus":0.2134837990183883,"score_gpt":0.3481269699009422,"score_spread":0.13464317088255393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385780698","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008722712,0.0017719964,0.9820004,0.0010965351,0.000063642714,0.0001557008,0.00022979421,0.003985376,0.0019737992],"genre_scores_gemma":[0.11912721,0.0012796231,0.8730247,0.0008610594,0.00010841284,0.00056191866,0.001727281,0.0009376808,0.002372096],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99439985,0.0033194898,0.00039530144,0.0008355748,0.0007906802,0.0002590369],"domain_scores_gemma":[0.9873017,0.007896454,0.00033121937,0.0022296838,0.0019383988,0.00030248525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011247408,0.0019932059,0.0015359919,0.0021171868,0.00097441644,0.0033521866,0.0054405443,0.0024233356,0.003930444],"category_scores_gemma":[0.031637,0.001845071,0.0012346596,0.0017995661,0.001361848,0.0067508956,0.004347054,0.005186656,0.003966081],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042345806,0.00048089566,0.0028293133,0.000718314,0.00036890703,0.00022298578,0.0007023633,0.2639135,0.007959372,0.038232088,0.01615703,0.6679917],"study_design_scores_gemma":[0.000102784004,0.00010407277,0.00022845507,0.0001306106,0.00006537911,0.00009816083,0.00018457959,0.93686825,0.0047490345,0.051162593,0.006273232,0.000032823056],"about_ca_topic_score_codex":0.0106134135,"about_ca_topic_score_gemma":0.025082622,"teacher_disagreement_score":0.011247408,"about_ca_system_score_codex":0.001967122,"about_ca_system_score_gemma":0.0031254026,"threshold_uncertainty_score":0.059482694},"labels":[],"label_agreement":null},{"id":"W4385800295","doi":"10.1007/s40747-023-01192-3","title":"An efficient long-text semantic retrieval approach via utilizing presentation learning on short-text","year":2023,"lang":"en","type":"article","venue":"Complex & Intelligent Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Science Foundation of Zhejiang Province; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Information retrieval; Relevance (law); Ranking (information retrieval); Document retrieval; Query expansion; Matching (statistics); Relevance feedback; Text retrieval; Artificial intelligence; Natural language processing; Image retrieval","score_opus":0.11265687229154413,"score_gpt":0.3291043899229319,"score_spread":0.21644751763138775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385800295","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058780447,0.0018076287,0.92730963,0.00047296006,0.00021704631,0.0002656339,0.00045521572,0.00651994,0.0041715163],"genre_scores_gemma":[0.72733194,0.0013097523,0.25202748,0.0004880265,0.00043155433,0.00031123322,0.0023314194,0.00031565846,0.015452913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999196,0.00014679287,0.00006364967,0.00019781681,0.00030970693,0.00008612095],"domain_scores_gemma":[0.99916387,0.0002257085,0.00007525623,0.00015345521,0.0003258197,0.000055938483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008421443,0.0009449179,0.0011281715,0.0020373196,0.00052554026,0.001049687,0.0015762758,0.0010046666,0.0040500406],"category_scores_gemma":[0.0021879172,0.00024208784,0.000893152,0.0018869191,0.0004637922,0.0039408146,0.0010598358,0.0008488719,0.0023533595],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065145676,0.00046057164,0.0017106911,0.00037997364,0.00008608079,0.00026756036,0.00015945539,0.08463809,0.049970716,0.009117963,0.01406028,0.83849716],"study_design_scores_gemma":[0.000046499477,0.00024677406,0.00056160137,0.000010385732,0.00004712097,0.00020597501,0.00006144784,0.97615063,0.012963318,0.0053675706,0.0043051606,0.000033652243],"about_ca_topic_score_codex":0.0035146018,"about_ca_topic_score_gemma":0.0029546898,"teacher_disagreement_score":0.0040500406,"about_ca_system_score_codex":0.00069825666,"about_ca_system_score_gemma":0.0011798694,"threshold_uncertainty_score":0.013548732},"labels":[],"label_agreement":null},{"id":"W4385967155","doi":"10.2196/48780","title":"Anki Tagger: A Generative AI Tool for Aligning Third-Party Resources to Preclinical Curriculum","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Education","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Curriculum; Generative grammar; Computer science; Artificial intelligence; Natural language processing; Psychology; Pedagogy","score_opus":0.029413606894392093,"score_gpt":0.3854807537735736,"score_spread":0.3560671468791815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385967155","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064603626,0.00029708527,0.8572825,0.0004087054,0.00015407437,0.00035548024,0.009372828,0.120287955,0.005381015],"genre_scores_gemma":[0.12164514,0.00046616668,0.8341335,0.00067005225,0.0001080013,0.0011149533,0.02430117,0.008888047,0.00867301],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99898976,0.00037383256,0.000088573965,0.00026392058,0.0002194199,0.00006446542],"domain_scores_gemma":[0.99512637,0.0038209946,0.00017359579,0.00046374265,0.00027645583,0.00013886021],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022566423,0.001765051,0.0007209126,0.0034613637,0.0007390801,0.0019670697,0.0020203767,0.001243734,0.020614643],"category_scores_gemma":[0.008634609,0.0010120773,0.0017206807,0.0020051962,0.0005628952,0.0022417884,0.0022424422,0.002144455,0.01047003],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072322227,0.00041208597,0.008440015,0.0016862749,0.00041065324,0.00065262086,0.0021719162,0.053054407,0.030574633,0.037059132,0.17413041,0.69068474],"study_design_scores_gemma":[0.00015482777,0.00015983061,0.0020525064,0.00017041778,0.00019961211,0.00042844986,0.00047982324,0.7398699,0.034044154,0.051658917,0.17063609,0.00014549385],"about_ca_topic_score_codex":0.008288576,"about_ca_topic_score_gemma":0.018083157,"teacher_disagreement_score":0.020614643,"about_ca_system_score_codex":0.0012628116,"about_ca_system_score_gemma":0.001984754,"threshold_uncertainty_score":0.06896287},"labels":[],"label_agreement":null},{"id":"W4385982296","doi":"10.1007/978-3-031-41682-8_7","title":"QuOTeS: Query-Oriented Technical Summarization","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval; Multi-document summarization; Usability; Upload; Query expansion; Query language; Web search query; World Wide Web; Search engine; Human–computer interaction","score_opus":0.02214431214374753,"score_gpt":0.25359150128722574,"score_spread":0.2314471891434782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385982296","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033039155,0.0019280503,0.90775895,0.0032945266,0.0013945973,0.00075346226,0.0077793263,0.0318636,0.041923538],"genre_scores_gemma":[0.058446832,0.0024801015,0.8015348,0.0013245284,0.0015591161,0.0007666762,0.033416122,0.010978721,0.08949311],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962734,0.0010700073,0.0004859251,0.0005001838,0.0014959377,0.00017452819],"domain_scores_gemma":[0.9930327,0.0025087257,0.0003192859,0.0013629433,0.0025891007,0.00018729421],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029483205,0.0017752141,0.0011904286,0.004165807,0.0015329578,0.005038812,0.0025643276,0.0012579892,0.06925625],"category_scores_gemma":[0.013598035,0.0009043653,0.0014107098,0.004842704,0.0007549986,0.006184822,0.003808814,0.001824784,0.046686914],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002316735,0.0000933168,0.0003870217,0.0013684937,0.00008510507,0.0002936699,0.0014262038,0.0026311825,0.013329621,0.08097877,0.36300403,0.5361709],"study_design_scores_gemma":[0.00006671443,0.00009434397,0.00048277536,0.00033379984,0.00013777814,0.00041400295,0.0010338675,0.039066862,0.019258142,0.10670081,0.8323121,0.00009872399],"about_ca_topic_score_codex":0.0015120346,"about_ca_topic_score_gemma":0.0018655699,"teacher_disagreement_score":0.06925625,"about_ca_system_score_codex":0.0007963409,"about_ca_system_score_gemma":0.0013069871,"threshold_uncertainty_score":0.23168522},"labels":[],"label_agreement":null},{"id":"W4386004792","doi":"10.1007/978-3-031-33261-6_8","title":"Topic Modelling for Automatically Identification of Relevant Concepts Discussed in Academic Documents","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Latent Dirichlet allocation; Topic model; Computer science; Identification (biology); Visualization; Coherence (philosophical gambling strategy); Data science; Information retrieval; Natural language processing; Artificial intelligence","score_opus":0.03895559091644764,"score_gpt":0.2946035290801671,"score_spread":0.2556479381637195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386004792","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07111449,0.011369614,0.88026613,0.0011457251,0.0007110201,0.000786229,0.01064433,0.016743382,0.007219179],"genre_scores_gemma":[0.41048032,0.0046944777,0.5406559,0.00037753434,0.0008724336,0.0013884435,0.030735169,0.001385836,0.009409954],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99771714,0.0007502048,0.00024096709,0.00056619843,0.0005039726,0.00022159966],"domain_scores_gemma":[0.9952179,0.0033296824,0.00028428494,0.00025777673,0.0007179354,0.00019245148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002449191,0.0014426433,0.0011584741,0.007861168,0.001183715,0.0027939954,0.0012859649,0.0017653458,0.004618056],"category_scores_gemma":[0.007787387,0.00056352734,0.0021922134,0.0057984237,0.00046956318,0.0033533496,0.0018880148,0.002009535,0.005081669],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013053406,0.00049206015,0.0077083176,0.0020522124,0.00039459034,0.00051540666,0.0019290127,0.012142446,0.044606812,0.010575921,0.065776765,0.8525011],"study_design_scores_gemma":[0.00021513886,0.0004081824,0.0147158895,0.00045541165,0.00085103675,0.0015654918,0.0016757182,0.8276526,0.03913978,0.03440118,0.07873441,0.00018521931],"about_ca_topic_score_codex":0.005109023,"about_ca_topic_score_gemma":0.005712027,"teacher_disagreement_score":0.007861168,"about_ca_system_score_codex":0.0011110911,"about_ca_system_score_gemma":0.0016698899,"threshold_uncertainty_score":0.015448928},"labels":[],"label_agreement":null},{"id":"W4386005452","doi":"10.1007/978-3-031-37963-5_80","title":"Attention Is not Always What You Need: Towards Efficient Classification of Domain-Specific Text","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada); Western University","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Classifier (UML); Jargon; Language model; Support vector machine; Question answering; Task (project management); Machine learning; Vectorization (mathematics); Domain (mathematical analysis); Linguistics; Mathematics","score_opus":0.04058312260523922,"score_gpt":0.23895681332880278,"score_spread":0.19837369072356356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386005452","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08417981,0.010367074,0.8715907,0.003936022,0.0009396836,0.00035114613,0.005564691,0.0132487,0.009822249],"genre_scores_gemma":[0.30140704,0.0055287965,0.63248545,0.0012388476,0.0016067411,0.00047694283,0.02180364,0.0016837645,0.03376881],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99907315,0.0002275718,0.00006347649,0.00029238334,0.00023073597,0.000112646536],"domain_scores_gemma":[0.9967321,0.0018175066,0.00018860017,0.00036438255,0.00070250646,0.0001949304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013461815,0.0011134598,0.0013814607,0.002766492,0.000791799,0.002302853,0.00186014,0.0015792766,0.004552521],"category_scores_gemma":[0.0052957977,0.0004990479,0.0010302514,0.003206602,0.0005667552,0.004251053,0.0019509122,0.0020681098,0.007832824],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028267707,0.00014239103,0.0017325303,0.0002897561,0.00006145921,0.000097226235,0.00037083332,0.0048614466,0.01387649,0.004029579,0.06679635,0.9074591],"study_design_scores_gemma":[0.0000776641,0.00024495088,0.005849003,0.00018326458,0.00019636402,0.00061083736,0.0011302346,0.8339569,0.025328225,0.06403249,0.0683141,0.00007596761],"about_ca_topic_score_codex":0.0035547279,"about_ca_topic_score_gemma":0.00538633,"teacher_disagreement_score":0.004552521,"about_ca_system_score_codex":0.0006822294,"about_ca_system_score_gemma":0.001129465,"threshold_uncertainty_score":0.015229702},"labels":[],"label_agreement":null},{"id":"W4386043297","doi":"10.48550/arxiv.2308.09124","title":"Linearity of Relation Decoding in Transformer Language Models","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation; Open Philanthropy Project","keywords":"Transformer; Linearity; Computer science; Decoding methods; Representation (politics); Computation; Relation (database); Variety (cybernetics); Artificial intelligence; Natural language processing; Theoretical computer science; Algorithm; Data mining; Electronic engineering","score_opus":0.15389028757315595,"score_gpt":0.21441950235854434,"score_spread":0.06052921478538839,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386043297","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052289948,0.00015421346,0.93733895,0.0006021623,0.000025900172,0.000047258873,0.00036551876,0.0013711001,0.0078049726],"genre_scores_gemma":[0.85734516,0.00028961952,0.13301374,0.0004105691,0.000078722886,0.00017191116,0.0012525108,0.0007820965,0.0066556768],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972573,0.0011902158,0.0001580894,0.00074235024,0.00041309817,0.00023892721],"domain_scores_gemma":[0.99003625,0.007323484,0.00045211238,0.0014991723,0.00052251475,0.00016650591],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024185956,0.000629329,0.0007546476,0.0009304287,0.00062461133,0.0031203479,0.0013719964,0.0008619099,0.0050506997],"category_scores_gemma":[0.019363973,0.00077723525,0.0012543787,0.0008313852,0.0022564318,0.00822147,0.002552955,0.0027466423,0.0020983196],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023860279,0.00007350083,0.0021285894,0.00020250573,0.0000724774,0.00029128918,0.00222915,0.1198743,0.0066338563,0.77848476,0.0030895656,0.08668137],"study_design_scores_gemma":[0.000013840047,0.00003184044,0.00022164464,0.00002070463,0.000022434831,0.00008312349,0.00013155463,0.3859794,0.0026604298,0.6094799,0.0013314891,0.000023631454],"about_ca_topic_score_codex":0.00517139,"about_ca_topic_score_gemma":0.004993743,"teacher_disagreement_score":0.00517139,"about_ca_system_score_codex":0.00180369,"about_ca_system_score_gemma":0.0012934315,"threshold_uncertainty_score":0.016896248},"labels":[],"label_agreement":null},{"id":"W4386074468","doi":"10.11159/cist23.152","title":"Historical-Domain Pre-trained Language Model for Historical Extractive Text Summarization","year":2023,"lang":"en","type":"article","venue":"Proceedings of the World Congress on Electrical Engineering and Computer Systems and Science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Natural language processing; Artificial intelligence; Domain (mathematical analysis); Language model; Information retrieval; Mathematics","score_opus":0.01208999847627023,"score_gpt":0.2206911119105644,"score_spread":0.20860111343429416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386074468","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019415542,0.0024277903,0.9615477,0.0004303958,0.00028914,0.00016986158,0.002483018,0.011202615,0.0020340737],"genre_scores_gemma":[0.29582492,0.0024440803,0.6535303,0.0006122202,0.0007726137,0.0008951583,0.028712265,0.0013851401,0.015823193],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9992193,0.00021301127,0.00006794667,0.00027941092,0.00014673555,0.00007366522],"domain_scores_gemma":[0.99852943,0.0005329427,0.00016340536,0.00018079304,0.0005398981,0.000053389038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011790033,0.0018046862,0.0010322885,0.0020825586,0.00049318594,0.00095046195,0.0014397592,0.0008629917,0.0029342473],"category_scores_gemma":[0.0034101754,0.00045602606,0.0012190064,0.0016234996,0.00037251355,0.0022973192,0.00085954374,0.0018756134,0.00462079],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040813533,0.00026268297,0.0015103248,0.00064987753,0.00021303797,0.00034379668,0.0004005765,0.104467295,0.049589507,0.004435744,0.02986633,0.8078527],"study_design_scores_gemma":[0.00006419387,0.00029490565,0.0012018593,0.000064056694,0.00015950348,0.00021502959,0.00024831507,0.9367129,0.030319393,0.0062473784,0.024407415,0.00006499319],"about_ca_topic_score_codex":0.0039583505,"about_ca_topic_score_gemma":0.007815027,"teacher_disagreement_score":0.0039583505,"about_ca_system_score_codex":0.00066551834,"about_ca_system_score_gemma":0.0013815291,"threshold_uncertainty_score":0.0098160505},"labels":[],"label_agreement":null},{"id":"W4386074631","doi":"10.11159/cist23.117","title":"Fine-Tuned PEGASUS: Exploring the Performance of the Transformer-Based Model on a Diverse Text Summarization Dataset","year":2023,"lang":"en","type":"article","venue":"Proceedings of the World Congress on Electrical Engineering and Computer Systems and Science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Transformer; Computer science; Information retrieval; Natural language processing; Artificial intelligence; Engineering; Electrical engineering; Voltage","score_opus":0.023889937776977347,"score_gpt":0.21235548674357702,"score_spread":0.18846554896659967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386074631","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75625265,0.015048156,0.14770038,0.0048584263,0.0015001765,0.0007324526,0.0289983,0.032398127,0.012511377],"genre_scores_gemma":[0.7841818,0.0016364125,0.12264046,0.0010805943,0.00038855214,0.0004035526,0.08028517,0.000995398,0.008388101],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990938,0.0003380651,0.00007405667,0.0003083165,0.00010737504,0.0000783335],"domain_scores_gemma":[0.9974039,0.0014276436,0.00014532817,0.00041002463,0.0004484943,0.00016454885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024515907,0.0017269003,0.00094819436,0.0019235665,0.0006534885,0.0011566085,0.001670721,0.0014623536,0.0019454336],"category_scores_gemma":[0.008518879,0.00021266784,0.0010489448,0.0013402074,0.0004958849,0.0025444138,0.0009995245,0.0018734292,0.0016990418],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024380446,0.0013550312,0.010491101,0.0015140012,0.00071238587,0.00043122782,0.0005698777,0.282599,0.015957953,0.0029652596,0.10933181,0.5716343],"study_design_scores_gemma":[0.00021092057,0.0008162597,0.0033033136,0.000077264376,0.00016356306,0.00013491824,0.00036205043,0.9691431,0.011196882,0.0040146285,0.010523494,0.000053531567],"about_ca_topic_score_codex":0.010758126,"about_ca_topic_score_gemma":0.01759175,"teacher_disagreement_score":0.010758126,"about_ca_system_score_codex":0.0013579606,"about_ca_system_score_gemma":0.0009386727,"threshold_uncertainty_score":0.021391034},"labels":[],"label_agreement":null},{"id":"W4386099755","doi":"10.1007/s00778-023-00809-w","title":"xDBTagger: explainable natural language interface to databases using keyword mappings and schema graph","year":2023,"lang":"en","type":"article","venue":"The VLDB Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Pipeline (software); Schema (genetic algorithms); SQL; Artificial intelligence; Natural language; Natural language processing; Information retrieval; Database; Programming language","score_opus":0.04688202870401947,"score_gpt":0.31181788788836495,"score_spread":0.2649358591843455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386099755","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005418815,0.0003288142,0.46128872,0.00055974774,0.00021545125,0.00034287656,0.049044024,0.47732016,0.0054812585],"genre_scores_gemma":[0.12501815,0.0014055549,0.5453085,0.0025686892,0.00018167049,0.0025089723,0.16805433,0.13125004,0.023704149],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99902844,0.000205286,0.00020204418,0.00024470768,0.00024706667,0.00007230552],"domain_scores_gemma":[0.99599683,0.0027886636,0.00018103547,0.00057045935,0.0003294005,0.00013359264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018661096,0.0019503069,0.0012829052,0.0020871141,0.00050222466,0.0028931885,0.0021786408,0.0017725923,0.052730847],"category_scores_gemma":[0.008149871,0.0013758439,0.001692777,0.0015415685,0.0005316115,0.004879174,0.0042923857,0.0017493347,0.01615648],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020332774,0.0002794673,0.005830723,0.004546542,0.00043067988,0.0019030588,0.0030390527,0.010154159,0.02555507,0.073255114,0.63228714,0.24068576],"study_design_scores_gemma":[0.0008741671,0.00021130366,0.0023503702,0.00093087717,0.0002588569,0.0012129932,0.00065674674,0.1311364,0.047231775,0.07378408,0.74107265,0.00027977867],"about_ca_topic_score_codex":0.0060851797,"about_ca_topic_score_gemma":0.0072100814,"teacher_disagreement_score":0.052730847,"about_ca_system_score_codex":0.0013378756,"about_ca_system_score_gemma":0.001237869,"threshold_uncertainty_score":0.17640227},"labels":[],"label_agreement":null},{"id":"W4386118460","doi":"10.3758/s13423-023-02360-9","title":"The prod eff: Partially producing items moderates the production effect","year":2023,"lang":"en","type":"article","venue":"Psychonomic Bulletin & Review","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Production (economics); Encoding (memory); Psychology; Context (archaeology); Feature (linguistics); Cognitive psychology; Linguistics; Microeconomics; Economics","score_opus":0.023895807239398513,"score_gpt":0.27752604261070474,"score_spread":0.25363023537130625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386118460","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9648495,0.002414624,0.0053331424,0.0013065987,0.00028237997,0.00015285822,0.0012313355,0.00028623085,0.024143254],"genre_scores_gemma":[0.98544115,0.0008903316,0.004533757,0.0007036546,0.00024332118,0.00014239558,0.0011330837,0.0003292845,0.0065830564],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982508,0.0007934462,0.00009373163,0.0004498451,0.00031286158,0.00009932224],"domain_scores_gemma":[0.94698167,0.04120819,0.003948103,0.0045719766,0.0014352911,0.0018548213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056844866,0.0011106684,0.0011301226,0.00048738488,0.00060602836,0.0024079368,0.0009646487,0.0017216476,0.026802875],"category_scores_gemma":[0.045913227,0.0010791878,0.00078962516,0.00033533573,0.001172412,0.0029325164,0.001586987,0.0024789032,0.0024285817],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.05564417,0.0067510037,0.3331109,0.005053142,0.004067837,0.001476489,0.0026962562,0.0038007605,0.1504625,0.013883505,0.013752429,0.40930107],"study_design_scores_gemma":[0.002040224,0.0065114764,0.9497099,0.00028943844,0.0030931015,0.00077224587,0.00038315685,0.0022516563,0.012187568,0.014535147,0.008115972,0.00011020324],"about_ca_topic_score_codex":0.0010761347,"about_ca_topic_score_gemma":0.0015040003,"teacher_disagreement_score":0.026802875,"about_ca_system_score_codex":0.0003863515,"about_ca_system_score_gemma":0.00061439374,"threshold_uncertainty_score":0.08966452},"labels":[],"label_agreement":null},{"id":"W4386207844","doi":"10.1109/iscc58397.2023.10217835","title":"MFG-R: Chinese Text Matching with Multi-Information Fusion Graph Embedding and Residual Connections","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"PricewaterhouseCoopers (Canada)","funders":"Natural Science Foundation of Heilongjiang Province","keywords":"Computer science; Artificial intelligence; Natural language processing; Residual; Word embedding; Graph; Matching (statistics); Text graph; Feature extraction; Embedding; Theoretical computer science; Mathematics; Algorithm","score_opus":0.015871298344720133,"score_gpt":0.2672689107356912,"score_spread":0.2513976123909711,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386207844","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05729087,0.0007796604,0.92590207,0.00038893268,0.00010350876,0.00027384816,0.0014659328,0.010258324,0.003536829],"genre_scores_gemma":[0.52701503,0.0005214767,0.45072073,0.0005030754,0.00011717364,0.0005278069,0.006895467,0.00076391327,0.012935389],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934155,0.000112750975,0.000037149886,0.00030022085,0.00013655912,0.00007174253],"domain_scores_gemma":[0.99961036,0.00008824945,0.000052863415,0.000117966316,0.00010064255,0.000029801478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00064945617,0.0011053176,0.00092472346,0.0018455434,0.00054394745,0.00066632946,0.0019792947,0.0012721295,0.0031206447],"category_scores_gemma":[0.0019553988,0.00033533224,0.0015042659,0.0017636586,0.00064564834,0.0029912968,0.0013983839,0.0008586506,0.0016153775],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004766693,0.00031234126,0.0044895397,0.00041882464,0.00023727748,0.00042460466,0.00039151983,0.14017744,0.028369548,0.026837742,0.021821847,0.77604264],"study_design_scores_gemma":[0.00003273393,0.000077429686,0.0010039426,0.000009864699,0.00003713215,0.00013193209,0.00004254519,0.9769483,0.0057731266,0.011864409,0.004053574,0.000025005207],"about_ca_topic_score_codex":0.016258566,"about_ca_topic_score_gemma":0.015222915,"teacher_disagreement_score":0.016258566,"about_ca_system_score_codex":0.00092608563,"about_ca_system_score_gemma":0.001081489,"threshold_uncertainty_score":0.03232789},"labels":[],"label_agreement":null},{"id":"W4386225280","doi":"10.32920/24043233.v1","title":"Assessment of Fine-Tuned GPT2 for Lyric Generation","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Lyrics; Computer science; Task (project management); Field (mathematics); Linguistics; Artificial intelligence; Art; Natural language processing; Natural (archaeology); Literature; History; Philosophy; Mathematics; Engineering","score_opus":0.15801961655639046,"score_gpt":0.36634761247331815,"score_spread":0.2083279959169277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386225280","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7539157,0.006580575,0.1905414,0.0028295787,0.0011193068,0.00054432615,0.0034983312,0.020205326,0.020765396],"genre_scores_gemma":[0.96194726,0.00038384058,0.030173928,0.00039668492,0.000084879466,0.00017633486,0.0032073269,0.00057360047,0.003056302],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921477,0.00030794094,0.000040670402,0.00023667325,0.00011754239,0.000082429586],"domain_scores_gemma":[0.99655765,0.0022513554,0.00013708705,0.0003848674,0.00043810575,0.00023089083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002947083,0.0015220455,0.0008105144,0.00078427646,0.00042552932,0.001374936,0.0015304602,0.0023185308,0.0028149998],"category_scores_gemma":[0.010200188,0.0005402482,0.00079387816,0.0005288762,0.000443664,0.0015018114,0.0010207123,0.002372361,0.001590704],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009878075,0.0005185351,0.005111765,0.0002979814,0.0003280125,0.00013887971,0.00009923251,0.8761241,0.0062059993,0.0011268965,0.0071919137,0.101868875],"study_design_scores_gemma":[0.00006649412,0.00015715795,0.00068798574,0.000017766064,0.000036145535,0.000024928182,0.00001878338,0.9951924,0.0024633298,0.0004782552,0.00084418716,0.0000125601655],"about_ca_topic_score_codex":0.01558601,"about_ca_topic_score_gemma":0.012710752,"teacher_disagreement_score":0.01558601,"about_ca_system_score_codex":0.0014300109,"about_ca_system_score_gemma":0.001169854,"threshold_uncertainty_score":0.0309906},"labels":[],"label_agreement":null},{"id":"W4386225388","doi":"10.32920/24043233","title":"Assessment of Fine-Tuned GPT2 for Lyric Generation","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Lyrics; Computer science; Task (project management); Field (mathematics); Natural language processing; Artificial intelligence; Linguistics; Language model; Natural (archaeology); Art; Speech recognition; Literature; Mathematics; History; Philosophy; Engineering","score_opus":0.15801961655639046,"score_gpt":0.36634761247331815,"score_spread":0.2083279959169277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386225388","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7539157,0.006580575,0.1905414,0.0028295787,0.0011193068,0.00054432615,0.0034983312,0.020205326,0.020765396],"genre_scores_gemma":[0.96194726,0.00038384058,0.030173928,0.00039668492,0.000084879466,0.00017633486,0.0032073269,0.00057360047,0.003056302],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921477,0.00030794094,0.000040670402,0.00023667325,0.00011754239,0.000082429586],"domain_scores_gemma":[0.99655765,0.0022513554,0.00013708705,0.0003848674,0.00043810575,0.00023089083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002947083,0.0015220455,0.0008105144,0.00078427646,0.00042552932,0.001374936,0.0015304602,0.0023185308,0.0028149998],"category_scores_gemma":[0.010200188,0.0005402482,0.00079387816,0.0005288762,0.000443664,0.0015018114,0.0010207123,0.002372361,0.001590704],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009878075,0.0005185351,0.005111765,0.0002979814,0.0003280125,0.00013887971,0.00009923251,0.8761241,0.0062059993,0.0011268965,0.0071919137,0.101868875],"study_design_scores_gemma":[0.00006649412,0.00015715795,0.00068798574,0.000017766064,0.000036145535,0.000024928182,0.00001878338,0.9951924,0.0024633298,0.0004782552,0.00084418716,0.0000125601655],"about_ca_topic_score_codex":0.01558601,"about_ca_topic_score_gemma":0.012710752,"teacher_disagreement_score":0.01558601,"about_ca_system_score_codex":0.0014300109,"about_ca_system_score_gemma":0.001169854,"threshold_uncertainty_score":0.0309906},"labels":[],"label_agreement":null},{"id":"W4386250690","doi":"10.24908/iqurcp16750","title":"Natural Language Processing of Radiology Reports: Predicting Metastatic Progression from Text Data","year":2023,"lang":"en","type":"article","venue":"Inquiry Queen s Undergraduate Research Conference Proceedings","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Readability; Context (archaeology); Natural language processing; Artificial intelligence; Sentence; Information retrieval; Unified Medical Language System; SNOMED CT; Radiology; Medical physics; Medicine; Terminology; Linguistics; Programming language","score_opus":0.16478113267487837,"score_gpt":0.4224478813233055,"score_spread":0.25766674864842715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386250690","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84961736,0.002633081,0.10797223,0.0026175382,0.00055545184,0.0008918755,0.020299673,0.011139257,0.004273549],"genre_scores_gemma":[0.81068933,0.0007423905,0.14082772,0.0004256661,0.00018506474,0.0004886009,0.043526877,0.00017787202,0.0029364745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99869615,0.00057543267,0.00013136388,0.0003497424,0.0001746299,0.00007268814],"domain_scores_gemma":[0.99132943,0.0067230943,0.00044571693,0.00045607102,0.00085594016,0.00018968423],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00277221,0.0012174932,0.00040778998,0.0016900628,0.00033171225,0.0012054254,0.0011201319,0.0011780878,0.0018957052],"category_scores_gemma":[0.01243254,0.00028336066,0.0009916646,0.0010386128,0.00028947723,0.0017981259,0.0006054441,0.0010992609,0.0017246742],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003099593,0.0021836373,0.067030184,0.001539297,0.00038358328,0.0024916134,0.0018720919,0.11041706,0.035243098,0.0019134221,0.036865253,0.73696125],"study_design_scores_gemma":[0.00015656347,0.00089119974,0.028407397,0.00012740993,0.00018454564,0.0007020844,0.00096673664,0.924234,0.026137024,0.0039548804,0.014146341,0.000091827525],"about_ca_topic_score_codex":0.0074693062,"about_ca_topic_score_gemma":0.009940336,"teacher_disagreement_score":0.0074693062,"about_ca_system_score_codex":0.0009307128,"about_ca_system_score_gemma":0.0008475722,"threshold_uncertainty_score":0.01485163},"labels":[],"label_agreement":null},{"id":"W4386256573","doi":"10.3390/stats6030056","title":"Investigating Self-Rationalizing Models for Commonsense Reasoning","year":2023,"lang":"en","type":"article","venue":"Stats","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Government of Canada","keywords":"Computer science; Generative grammar; Leverage (statistics); Commonsense reasoning; Language model; Transformer; Artificial intelligence; Natural language processing; Natural language understanding; Generative model; Natural language; Representation (politics)","score_opus":0.09268685221712009,"score_gpt":0.31079395343024085,"score_spread":0.21810710121312077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386256573","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.294369,0.001272553,0.6840947,0.004038097,0.00016466354,0.0002582987,0.0024544627,0.0053078365,0.008040409],"genre_scores_gemma":[0.900293,0.00024699885,0.09315348,0.0004929931,0.000081012404,0.00016008716,0.0034013703,0.00033397865,0.0018370413],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99791604,0.0010456747,0.000099609315,0.0005786843,0.00023228428,0.0001276417],"domain_scores_gemma":[0.97573084,0.020453798,0.00066505396,0.0021491535,0.00061691925,0.00038421017],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00615972,0.0007738612,0.0006055295,0.0011891384,0.0005176503,0.0024659506,0.0016282821,0.0014080462,0.004106373],"category_scores_gemma":[0.0288867,0.0004254569,0.0013951875,0.0007126151,0.0014533334,0.005387784,0.0021300723,0.003247522,0.0010540672],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077332795,0.00064006454,0.020503916,0.0008158574,0.00040836562,0.00037907672,0.0020957969,0.5277526,0.0075061633,0.2320988,0.014196371,0.19282976],"study_design_scores_gemma":[0.00003780429,0.000032583775,0.00042283098,0.000022368396,0.000020713363,0.000038810493,0.00007910766,0.90343636,0.0013726945,0.093133315,0.0013924304,0.000010963332],"about_ca_topic_score_codex":0.0027755457,"about_ca_topic_score_gemma":0.0051232427,"teacher_disagreement_score":0.00615972,"about_ca_system_score_codex":0.0016978728,"about_ca_system_score_gemma":0.0012755546,"threshold_uncertainty_score":0.032576144},"labels":[],"label_agreement":null},{"id":"W4386284106","doi":"10.18280/ts.400406","title":"Multimodal Deep Learning Framework for Book Recommendations: Harnessing Image Processing with VGG16 and Textual Analysis via LSTM-Enhanced Word2Vec","year":2023,"lang":"en","type":"article","venue":"Traitement du signal","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Word2vec; Computer science; Artificial intelligence; Deep learning; Sentiment analysis; Image (mathematics); Natural language processing; Pattern recognition (psychology); Embedding","score_opus":0.017688035973752746,"score_gpt":0.27830852262750666,"score_spread":0.2606204866537539,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386284106","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057882678,0.0010689751,0.931652,0.00052754313,0.00014064908,0.00006892971,0.0006448808,0.0042349403,0.0037793156],"genre_scores_gemma":[0.6979011,0.0008406215,0.28673676,0.000382576,0.00012791528,0.00013966828,0.0018796368,0.000256664,0.011735004],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998822,0.000021478103,0.000005559582,0.000036725385,0.00003041301,0.00002352566],"domain_scores_gemma":[0.99987006,0.00004033391,0.000011368927,0.000017808678,0.000048693997,0.000011736046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00027054673,0.00078817195,0.00042902338,0.00061056466,0.00017822508,0.0005400551,0.0007837083,0.0006851713,0.0022981567],"category_scores_gemma":[0.0008165132,0.00024264175,0.00045783498,0.000773302,0.0001976315,0.00083726185,0.00049973745,0.0009888079,0.0014617061],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019548523,0.00017806281,0.0016995935,0.00010878742,0.00013879493,0.00012398453,0.00011287726,0.23965727,0.025189769,0.0045732143,0.008309375,0.71971273],"study_design_scores_gemma":[0.0000043316572,0.000033186658,0.00024792794,0.000006451794,0.000012089674,0.000017376919,0.000013068901,0.99389946,0.0032238674,0.0016509654,0.0008853789,0.000005821832],"about_ca_topic_score_codex":0.011571537,"about_ca_topic_score_gemma":0.023010863,"teacher_disagreement_score":0.011571537,"about_ca_system_score_codex":0.0005438,"about_ca_system_score_gemma":0.00057831657,"threshold_uncertainty_score":0.023008347},"labels":[],"label_agreement":null},{"id":"W4386332795","doi":"10.3390/aerospace10090770","title":"Examining the Potential of Generative Language Models for Aviation Safety Analysis: Case Study and Insights Using the Aviation Safety Reporting System (ASRS)","year":2023,"lang":"en","type":"article","venue":"Aerospace","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"General Fusion (Canada)","funders":"","keywords":"Aviation; Aviation safety; Computer science; Context (archaeology); Process (computing); Generative grammar; Risk analysis (engineering); Engineering; Artificial intelligence; Business","score_opus":0.06448011856139108,"score_gpt":0.30374454778840826,"score_spread":0.23926442922701718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386332795","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60596764,0.0014278479,0.3743319,0.0040745754,0.00021481344,0.00083057024,0.0027198675,0.0036860676,0.006746795],"genre_scores_gemma":[0.79526997,0.00037006076,0.19873947,0.0004012967,0.000051839586,0.0003889393,0.0032909743,0.00031883898,0.0011686092],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9880497,0.009542285,0.0003572811,0.0009233154,0.00091854733,0.00020886828],"domain_scores_gemma":[0.9304671,0.0617543,0.0017404337,0.0034440926,0.0021527056,0.00044129053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011518954,0.0013456432,0.00046639013,0.002531924,0.00077077997,0.002922602,0.0015299049,0.0015684814,0.0014954212],"category_scores_gemma":[0.048407253,0.00046697666,0.0013800748,0.0015271235,0.0013869697,0.0033442194,0.0024135157,0.0019044046,0.000624848],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018644449,0.0016630886,0.097416095,0.0029974072,0.0006192999,0.004243536,0.03239468,0.46476436,0.018762212,0.042422794,0.012977515,0.31987455],"study_design_scores_gemma":[0.00009403002,0.0005650721,0.0072672907,0.00036775626,0.00015428683,0.0007713002,0.0055986037,0.9401462,0.008687151,0.021139804,0.01507378,0.00013480359],"about_ca_topic_score_codex":0.008008091,"about_ca_topic_score_gemma":0.011195405,"teacher_disagreement_score":0.011518954,"about_ca_system_score_codex":0.0018378062,"about_ca_system_score_gemma":0.0016863173,"threshold_uncertainty_score":0.060918808},"labels":[],"label_agreement":null},{"id":"W4386422441","doi":"10.1007/s11227-023-05592-7","title":"SiMaLSTM-SNP: novel semantic relatedness learning model preserving both Siamese networks and membrane computing","year":2023,"lang":"en","type":"article","venue":"The Journal of Supercomputing","topic":"Topic Modeling","field":"Computer Science","cited_by":87,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Sentence; Task (project management); Semantic similarity; Deep learning; State (computer science); Algorithm","score_opus":0.03696003122406642,"score_gpt":0.26234672475837906,"score_spread":0.22538669353431262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386422441","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018669888,0.00068557396,0.9644739,0.0008291138,0.00039551264,0.0001145153,0.001048315,0.010628949,0.0031542748],"genre_scores_gemma":[0.38785386,0.0007006026,0.5880568,0.001274608,0.0004183352,0.0004242921,0.006011335,0.0013101951,0.013950006],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896026,0.00025203946,0.000049996685,0.00040060582,0.00023935552,0.0000976916],"domain_scores_gemma":[0.9993284,0.00017618385,0.00003541745,0.00022440968,0.00016866115,0.00006698859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011864902,0.0009462071,0.0015079597,0.0008257846,0.00084480294,0.0013826374,0.00290655,0.0021618642,0.0040053804],"category_scores_gemma":[0.0037599157,0.00046062915,0.001104736,0.0013123457,0.00066920597,0.0038923207,0.0029073071,0.0026009898,0.0036987613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000419944,0.000538682,0.0015027535,0.00022690541,0.00022322402,0.0001991215,0.0001352718,0.119575836,0.0142281335,0.036513306,0.05296925,0.7734676],"study_design_scores_gemma":[0.00002259244,0.000052736148,0.00015083037,0.000007276876,0.000021286025,0.00005861106,0.00001864183,0.96268886,0.0030106625,0.030060986,0.0038915102,0.000016001306],"about_ca_topic_score_codex":0.005529228,"about_ca_topic_score_gemma":0.009895451,"teacher_disagreement_score":0.005529228,"about_ca_system_score_codex":0.00076319557,"about_ca_system_score_gemma":0.00227484,"threshold_uncertainty_score":0.013399303},"labels":[],"label_agreement":null},{"id":"W4386429868","doi":"10.1007/978-981-99-5837-5_25","title":"Novel Topic Models for Parallel Topics Extraction from Multilingual Text","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Jaccard index; Topic model; Flexibility (engineering); Inference; Latent Dirichlet allocation; Prior probability; Artificial intelligence; Dirichlet distribution; Information retrieval; Natural language processing; Machine learning; Data mining; Pattern recognition (psychology); Mathematics; Boundary value problem; Statistics","score_opus":0.06478765766729977,"score_gpt":0.2963527200690191,"score_spread":0.2315650624017193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386429868","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011208198,0.0020933151,0.9744105,0.0003885935,0.00034223535,0.0001795109,0.0026007318,0.0070463675,0.0017307277],"genre_scores_gemma":[0.16954379,0.0025755398,0.7915421,0.00028116148,0.0009387488,0.0010556772,0.021530926,0.0023737634,0.0101583],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982686,0.00046155127,0.00018569251,0.0006161798,0.00032215,0.00014582959],"domain_scores_gemma":[0.99772114,0.0013069197,0.00014179581,0.000294787,0.0004322357,0.000103124716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020578867,0.0019800598,0.0014711308,0.003328155,0.0011810435,0.0027584422,0.0015179348,0.0015087884,0.005253593],"category_scores_gemma":[0.0050676353,0.0009126415,0.0026156837,0.0043064957,0.00042982612,0.005098628,0.00211137,0.0028133218,0.0065300427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009868513,0.00032289562,0.0024386344,0.0008275178,0.0005111322,0.00036119096,0.000997421,0.02456812,0.032758582,0.016941834,0.04903371,0.8702521],"study_design_scores_gemma":[0.0001340453,0.0001348451,0.0024421779,0.00011945134,0.000382198,0.00047768187,0.00045045197,0.89924455,0.019055927,0.03718197,0.04026871,0.000108048545],"about_ca_topic_score_codex":0.0049553066,"about_ca_topic_score_gemma":0.008145944,"teacher_disagreement_score":0.005253593,"about_ca_system_score_codex":0.00096700346,"about_ca_system_score_gemma":0.0016874494,"threshold_uncertainty_score":0.017574966},"labels":[],"label_agreement":null},{"id":"W4386471798","doi":"10.1186/s42400-023-00183-8","title":"Use of subword tokenization for domain generation algorithm classification","year":2023,"lang":"en","type":"article","venue":"Cybersecurity","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Lexical analysis; Computer science; Artificial intelligence; Word (group theory); Convolutional neural network; Feature (linguistics); Domain (mathematical analysis); Scheme (mathematics); Machine learning; Natural language processing; Mathematics","score_opus":0.11790842494516683,"score_gpt":0.2921493607814226,"score_spread":0.17424093583625577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386471798","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3487353,0.0010084534,0.63237816,0.0003784527,0.0004583313,0.0004396438,0.001268638,0.010852281,0.0044806735],"genre_scores_gemma":[0.8196653,0.00017384576,0.17425993,0.00009116919,0.00007812049,0.00015412123,0.0028501763,0.00022648434,0.0025009643],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99831855,0.00048684762,0.00020687826,0.0004629189,0.00033613748,0.00018881983],"domain_scores_gemma":[0.9962728,0.001215169,0.0004359384,0.0008558608,0.00097040395,0.00024978496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014629386,0.000741331,0.00082177453,0.0030389791,0.0006177109,0.0015290218,0.0012811482,0.00067566027,0.0023866198],"category_scores_gemma":[0.005815296,0.00016441214,0.00059868273,0.0020566706,0.0004426383,0.0022709095,0.0012072668,0.0010221428,0.00192539],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009717415,0.00032829837,0.025209425,0.00019089284,0.00011377684,0.000208517,0.00019613333,0.025539763,0.026359221,0.004999776,0.007981365,0.90790105],"study_design_scores_gemma":[0.000032525284,0.00014284156,0.004218621,0.000016307798,0.000054331187,0.00016453197,0.00010747105,0.956163,0.031581175,0.0032737555,0.0042108013,0.000034666333],"about_ca_topic_score_codex":0.0033374846,"about_ca_topic_score_gemma":0.003324644,"teacher_disagreement_score":0.0033374846,"about_ca_system_score_codex":0.0008600656,"about_ca_system_score_gemma":0.0014230127,"threshold_uncertainty_score":0.007984042},"labels":[],"label_agreement":null},{"id":"W4386493132","doi":"10.1109/icnlp58431.2023.00057","title":"Context-aware Information Extraction from Multi-thread Business Conversations","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Systems, Applications & Products in Data Processing (Canada)","funders":"","keywords":"Conversation; Computer science; Sentence; Parsing; Thread (computing); Natural language processing; Information extraction; Context (archaeology); Artificial intelligence; Linguistics; Programming language","score_opus":0.048635577052965784,"score_gpt":0.27843268030035717,"score_spread":0.2297971032473914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386493132","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044011362,0.0011202114,0.9434716,0.0005300682,0.00016371958,0.00040644722,0.0018881112,0.004328622,0.0040797847],"genre_scores_gemma":[0.28831533,0.0009234003,0.70084023,0.00014606454,0.0002363403,0.00040701378,0.004738817,0.00043840831,0.003954301],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986351,0.00041586542,0.0001245814,0.00039092696,0.0002980162,0.00013559677],"domain_scores_gemma":[0.99738985,0.0013890814,0.00023633814,0.0002655245,0.0006175676,0.00010160628],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014076125,0.0014258927,0.0009391431,0.0034669966,0.0011779885,0.0019163925,0.00082197937,0.0011266189,0.0020226182],"category_scores_gemma":[0.0054632863,0.0006266748,0.001104555,0.0021257934,0.00035352286,0.0030710688,0.0020025955,0.0014585964,0.0021126652],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008600061,0.00032993758,0.007808796,0.0012230439,0.00018729325,0.0015175537,0.005237383,0.012375115,0.113876894,0.012225763,0.012672009,0.8316862],"study_design_scores_gemma":[0.00007933008,0.00036689904,0.018495271,0.0005043775,0.00045158673,0.0015475749,0.00622745,0.6782157,0.15222773,0.06457051,0.077075705,0.0002378744],"about_ca_topic_score_codex":0.0018876133,"about_ca_topic_score_gemma":0.0029538865,"teacher_disagreement_score":0.0034669966,"about_ca_system_score_codex":0.0005185575,"about_ca_system_score_gemma":0.0012384198,"threshold_uncertainty_score":0.0074442625},"labels":[],"label_agreement":null},{"id":"W4386517708","doi":"10.1145/3594536.3595176","title":"Summary of the Competition on Legal Information, Extraction/Entailment (COLIEE) 2023","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Task (project management); Statute; Computer science; Logical consequence; Component (thermodynamics); Competition (biology); Common law; Variety (cybernetics); Natural language processing; Information retrieval; Artificial intelligence; Law; Political science; Engineering","score_opus":0.020991304753814995,"score_gpt":0.2554400655134253,"score_spread":0.2344487607596103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386517708","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032943886,0.031648636,0.07931796,0.04184157,0.078782745,0.008184481,0.43264008,0.0696593,0.22498141],"genre_scores_gemma":[0.021166379,0.003952743,0.052334122,0.005330776,0.0051811365,0.003028242,0.7336256,0.01723052,0.15815051],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97481155,0.006308341,0.0015839918,0.0028677336,0.011865926,0.002562407],"domain_scores_gemma":[0.942813,0.008799182,0.0009060956,0.004837363,0.030977203,0.011667137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028701361,0.0064064823,0.004930552,0.008107696,0.005064487,0.014109938,0.006839206,0.0049750274,0.10073614],"category_scores_gemma":[0.04415308,0.0017105637,0.00399839,0.009579036,0.0013872525,0.007590778,0.008411582,0.005455698,0.099033184],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020762977,0.0001606225,0.00018442614,0.00030905838,0.00003817448,0.00003469772,0.00004203785,0.000529885,0.0005498698,0.000435917,0.9733251,0.024182688],"study_design_scores_gemma":[0.0005018945,0.000327685,0.0039485344,0.00030212177,0.00007635077,0.00018083012,0.00020861845,0.007876558,0.0030701475,0.0033379463,0.98001295,0.00015639604],"about_ca_topic_score_codex":0.048079856,"about_ca_topic_score_gemma":0.09656691,"teacher_disagreement_score":0.10073614,"about_ca_system_score_codex":0.007307955,"about_ca_system_score_gemma":0.012202481,"threshold_uncertainty_score":0.33699596},"labels":[],"label_agreement":null},{"id":"W4386558807","doi":"10.1109/tg.2023.3313121","title":"Leveraging the OPT Large Language Model for Sentiment Analysis of Game Reviews","year":2023,"lang":"en","type":"article","venue":"IEEE Transactions on Games","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Terminology; Sentiment analysis; Classifier (UML); Artificial intelligence; Natural language processing; Field (mathematics); Machine learning; Linguistics","score_opus":0.0512710427285866,"score_gpt":0.3127588787193115,"score_spread":0.2614878359907249,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386558807","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2156594,0.0024679897,0.7476312,0.002558081,0.0010013044,0.00092632486,0.0070878114,0.009610495,0.013057415],"genre_scores_gemma":[0.7524863,0.0007870889,0.22179745,0.00087758387,0.00050851004,0.0007419067,0.011195897,0.0003226373,0.0112826],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989272,0.0004023776,0.00010045845,0.00024176651,0.0002592077,0.0000688888],"domain_scores_gemma":[0.99826103,0.0007216015,0.00016552258,0.00012595009,0.0006589989,0.00006690345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018050022,0.00096082874,0.00058691023,0.0017989016,0.00041822638,0.0016615655,0.00075399963,0.00060314377,0.0016235312],"category_scores_gemma":[0.0046115154,0.0003130232,0.0011043067,0.0009904933,0.00022927421,0.0014829194,0.00062233105,0.0011727973,0.0027101457],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012174069,0.00092750887,0.036952917,0.0005636581,0.0006053409,0.0004567338,0.0006260479,0.086918734,0.02999174,0.008194376,0.051327605,0.782218],"study_design_scores_gemma":[0.000030901036,0.0001074587,0.004343082,0.000022165123,0.000059129092,0.00009586996,0.00007636411,0.9839792,0.0021539591,0.0038444363,0.0052615893,0.000025792497],"about_ca_topic_score_codex":0.006989785,"about_ca_topic_score_gemma":0.014508876,"teacher_disagreement_score":0.006989785,"about_ca_system_score_codex":0.0008661237,"about_ca_system_score_gemma":0.0011467563,"threshold_uncertainty_score":0.013898194},"labels":[],"label_agreement":null},{"id":"W4386566463","doi":"10.18653/v1/2023.findings-eacl.143","title":"Towards Fine-tuning Pre-trained Language Models with Integer Forward and Backward Propagation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Integer (computer science); Association (psychology); Natural language processing; Linguistics; Artificial intelligence; Programming language; Psychology; Philosophy","score_opus":0.021841055211201273,"score_gpt":0.2550813753600335,"score_spread":0.23324032014883223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566463","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052789204,0.005454825,0.89754516,0.0012048193,0.0009870986,0.00020193645,0.0012986901,0.034121174,0.006396963],"genre_scores_gemma":[0.4087979,0.00164774,0.56271327,0.0011922133,0.0005493326,0.00044747436,0.008161904,0.0038127257,0.012677448],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999042,0.0003025621,0.00007020682,0.0003144052,0.00012480367,0.00014593644],"domain_scores_gemma":[0.9969278,0.0020313314,0.000082089995,0.00025231685,0.0005706855,0.00013571189],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021653487,0.0025495528,0.0016811625,0.0014423514,0.0008462556,0.0024008877,0.002855107,0.0021963771,0.005341416],"category_scores_gemma":[0.005921771,0.0013421783,0.0015041499,0.0013181217,0.0005464986,0.0038218324,0.0019449873,0.004632712,0.006950638],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072448165,0.00053913967,0.0020009598,0.00043404129,0.0004458575,0.0002833871,0.00032389452,0.20874001,0.017809372,0.0048964205,0.041665025,0.72213745],"study_design_scores_gemma":[0.000044548113,0.000034356617,0.0001709894,0.000023056135,0.00005752795,0.000032283584,0.000044080458,0.9911221,0.002848871,0.0037340021,0.0018755745,0.000012485984],"about_ca_topic_score_codex":0.020132042,"about_ca_topic_score_gemma":0.038426466,"teacher_disagreement_score":0.020132042,"about_ca_system_score_codex":0.0011634097,"about_ca_system_score_gemma":0.0025098096,"threshold_uncertainty_score":0.040029705},"labels":[],"label_agreement":null},{"id":"W4386566488","doi":"10.18653/v1/2023.findings-eacl.83","title":"Large Language Models are few(1)-shot Table Reasoners","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Vector Institute","funders":"","keywords":"Table (database); Context (archaeology); Computer science; Shot (pellet); Code (set theory); Natural language processing; Artificial intelligence; Programming language; Database; Geography","score_opus":0.04316837089407304,"score_gpt":0.275202534745373,"score_spread":0.23203416385129993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566488","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03216017,0.0011099582,0.92502046,0.0014669261,0.00027255304,0.0002674234,0.002738631,0.030727467,0.006236422],"genre_scores_gemma":[0.29740208,0.0004442136,0.6852754,0.0012382737,0.00014787399,0.0003478703,0.0065690503,0.001730536,0.006844722],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99795014,0.00086021237,0.00012060383,0.0006307208,0.0003258155,0.000112419955],"domain_scores_gemma":[0.9933094,0.0047224294,0.00023145313,0.0011702402,0.00039997484,0.0001664878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031458074,0.0014891608,0.0009781884,0.0010028807,0.00067757186,0.0026840244,0.0025189333,0.0017413865,0.012039439],"category_scores_gemma":[0.016471243,0.0008747914,0.0020491993,0.0007336987,0.0009627004,0.0064613493,0.0021639708,0.003455807,0.005998746],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012324627,0.0004813036,0.005981897,0.0017818692,0.0005441452,0.0005243173,0.0016301278,0.1641717,0.019205347,0.06732272,0.07384965,0.6632746],"study_design_scores_gemma":[0.00011368739,0.00013847365,0.00050788245,0.00011197587,0.00009651441,0.00020588878,0.00031727177,0.875646,0.007850339,0.09568312,0.019279018,0.000049959737],"about_ca_topic_score_codex":0.0045111952,"about_ca_topic_score_gemma":0.0141590275,"teacher_disagreement_score":0.012039439,"about_ca_system_score_codex":0.0010994095,"about_ca_system_score_gemma":0.0016347178,"threshold_uncertainty_score":0.04027599},"labels":[],"label_agreement":null},{"id":"W4386566497","doi":"10.18653/v1/2023.findings-eacl.194","title":"Discourse Structure Extraction from Pre-Trained and Fine-Tuned Language Models in Dialogues","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Agence Nationale de la Recherche","keywords":"Computer science; Natural language processing; Artificial intelligence; Task (project management); Sentence; Exploit; Language model","score_opus":0.02511427457892696,"score_gpt":0.2887800994238192,"score_spread":0.26366582484489226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2067744,0.004611266,0.747918,0.001304848,0.0007057655,0.00037949928,0.006416289,0.024804188,0.0070857042],"genre_scores_gemma":[0.7675643,0.0010221275,0.20698516,0.0001826897,0.00031790356,0.00035132852,0.0151057765,0.0015327322,0.006938053],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99839383,0.0006091728,0.00009645115,0.0006032131,0.0001548674,0.00014254705],"domain_scores_gemma":[0.9961971,0.0027467855,0.00011579054,0.00025294427,0.00055807154,0.00012939464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018672495,0.0018013511,0.001183932,0.0019873127,0.00078107236,0.0021365301,0.0010705063,0.0016102447,0.0040314496],"category_scores_gemma":[0.0067698373,0.00068216084,0.0013879568,0.0011157374,0.00037732825,0.0022579795,0.001336362,0.0022558018,0.0048843557],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017730708,0.000405261,0.004667478,0.0011441426,0.00036709406,0.000523369,0.0019400524,0.04868875,0.10346221,0.0031749005,0.023023007,0.81083065],"study_design_scores_gemma":[0.00011514455,0.0002500736,0.0049393578,0.00013406297,0.00031370515,0.00020952697,0.0010095747,0.9193135,0.050771102,0.006949517,0.015917808,0.000076643424],"about_ca_topic_score_codex":0.0041010724,"about_ca_topic_score_gemma":0.004964642,"teacher_disagreement_score":0.0041010724,"about_ca_system_score_codex":0.00087058934,"about_ca_system_score_gemma":0.0012848899,"threshold_uncertainty_score":0.013486564},"labels":[],"label_agreement":null},{"id":"W4386566583","doi":"10.18653/v1/2023.findings-eacl.12","title":"Revisiting Intermediate Layer Distillation for Compressing Language Models: An Overfitting Perspective","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Overfitting; Computer science; Distillation; Code (set theory); Benchmark (surveying); Machine learning; Artificial intelligence; Consistency (knowledge bases); Simple (philosophy); Transformer; Layer (electronics); Language model; Natural language processing; Artificial neural network; Programming language; Engineering; Chemistry","score_opus":0.0662106754848642,"score_gpt":0.33617056808514506,"score_spread":0.26995989260028086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566583","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02387727,0.0007869448,0.9678189,0.00063740613,0.0001006305,0.00008730482,0.0004516497,0.0046449564,0.0015948866],"genre_scores_gemma":[0.40380418,0.0008486578,0.58163404,0.0011453422,0.00021801758,0.00034475612,0.0038314932,0.0013027609,0.0068706153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99884343,0.00038729724,0.000092338436,0.00027787563,0.00026655183,0.00013250201],"domain_scores_gemma":[0.99659497,0.0018770339,0.00015279699,0.0008732528,0.0003798345,0.0001221744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026203005,0.0016426728,0.0013748758,0.0012518896,0.0007569592,0.0021298048,0.0028355778,0.0014622491,0.0044526854],"category_scores_gemma":[0.014068878,0.00065212627,0.001337629,0.0017444403,0.0014229029,0.005099816,0.0039225663,0.00454647,0.0025046458],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041135465,0.00022781195,0.0026166937,0.00042113973,0.00018686581,0.00026312628,0.0003293504,0.315048,0.013306869,0.03222757,0.01150181,0.6234594],"study_design_scores_gemma":[0.000032096214,0.00006261216,0.00020508957,0.000030186764,0.000025112466,0.00007555249,0.000051986073,0.9682076,0.0066407905,0.02139561,0.0032489672,0.000024455056],"about_ca_topic_score_codex":0.005960368,"about_ca_topic_score_gemma":0.011638204,"teacher_disagreement_score":0.005960368,"about_ca_system_score_codex":0.0008820301,"about_ca_system_score_gemma":0.0027222033,"threshold_uncertainty_score":0.014895678},"labels":[],"label_agreement":null},{"id":"W4386566601","doi":"10.18653/v1/2023.findings-eacl.52","title":"Detecting Contextomized Quotes in News Headlines by Contrastive Learning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Institute for Basic Science; Korea Advanced Institute of Science and Technology; National Research Foundation of Korea; National Research Foundation","keywords":"Headline; Computer science; Context (archaeology); Credibility; Citation; Sign (mathematics); Matching (statistics); Embedding; Code (set theory); Domain (mathematical analysis); Appeal; Information retrieval; Natural language processing; Linguistics; Artificial intelligence; World Wide Web; History; Programming language; Political science; Mathematics; Philosophy; Epistemology","score_opus":0.027099880597303106,"score_gpt":0.27766790709469347,"score_spread":0.25056802649739035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566601","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67455316,0.017476687,0.16522904,0.003246803,0.0029897604,0.0010789678,0.083756804,0.02189754,0.02977132],"genre_scores_gemma":[0.7539558,0.0016554091,0.11947414,0.00081777835,0.0015302049,0.00068337197,0.10750658,0.0011527085,0.013224141],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99850523,0.00046032,0.000118380725,0.00054576737,0.000252652,0.00011758848],"domain_scores_gemma":[0.99354255,0.0037575404,0.0008015458,0.0007733094,0.00086752616,0.00025739922],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018678273,0.0013434957,0.0005477538,0.004264505,0.00079045893,0.0021549342,0.0010305439,0.0015706284,0.0032634404],"category_scores_gemma":[0.010450628,0.0003112514,0.0007254149,0.002319965,0.00074492604,0.0026642112,0.0017250667,0.0018771206,0.004382093],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022090604,0.0010660915,0.08263022,0.0029370154,0.00039404142,0.0017644686,0.004909512,0.013000407,0.050988335,0.008008389,0.26804206,0.5640504],"study_design_scores_gemma":[0.00051808846,0.0008775818,0.08752234,0.0010086163,0.0005585452,0.00326323,0.0069081355,0.48535842,0.06372583,0.039853882,0.3100815,0.000323748],"about_ca_topic_score_codex":0.0028373122,"about_ca_topic_score_gemma":0.0087119285,"teacher_disagreement_score":0.004264505,"about_ca_system_score_codex":0.0006092811,"about_ca_system_score_gemma":0.00062446436,"threshold_uncertainty_score":0.010917306},"labels":[],"label_agreement":null},{"id":"W4386566626","doi":"10.18653/v1/2023.eacl-main.88","title":"Policy-based Reinforcement Learning for Generalisation in Interactive Text-based Environments","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut National de la Recherche Scientifique","funders":"University of Cape Town; National Research Foundation","keywords":"Converse; Reinforcement learning; Computer science; Benchmark (surveying); Natural language; Artificial intelligence; Variety (cybernetics); Simple (philosophy); Question answering; Value (mathematics); Natural language understanding; Baseline (sea); Machine learning; Mathematics","score_opus":0.035474172160724,"score_gpt":0.2884375952942587,"score_spread":0.2529634231335347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566626","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06008566,0.00027331413,0.9338612,0.0003178331,0.000055007287,0.00013110666,0.000078309764,0.0028416202,0.0023560624],"genre_scores_gemma":[0.85633415,0.00013646229,0.14023247,0.00022906657,0.00004424822,0.0002497914,0.0002122564,0.00029214856,0.0022694143],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987463,0.00062355585,0.00005725381,0.0003065121,0.00015302097,0.00011342722],"domain_scores_gemma":[0.9942656,0.004519222,0.00024424997,0.00043365915,0.00034365253,0.00019358468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002773597,0.0010758701,0.0012315395,0.0005540589,0.00049539004,0.0011001198,0.002016339,0.0013260052,0.0024331189],"category_scores_gemma":[0.013387931,0.00047765393,0.00054690195,0.0003981184,0.0014585131,0.0023900112,0.0016475596,0.002136103,0.00062711333],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022522883,0.00021409962,0.0014154982,0.00015446731,0.00005992802,0.00015042457,0.00040543178,0.8791914,0.005983117,0.00897632,0.0015173174,0.101706766],"study_design_scores_gemma":[0.000018142677,0.000036215468,0.00009069669,0.000005824892,0.000005058023,0.000011275158,0.000019976462,0.99370414,0.00087026187,0.004892911,0.00033868826,0.000006775861],"about_ca_topic_score_codex":0.0051497957,"about_ca_topic_score_gemma":0.004898888,"teacher_disagreement_score":0.0051497957,"about_ca_system_score_codex":0.0011074804,"about_ca_system_score_gemma":0.00095830363,"threshold_uncertainty_score":0.014668345},"labels":[],"label_agreement":null},{"id":"W4386566650","doi":"10.18653/v1/2023.eacl-main.205","title":"Bridging the Gap Between BabelNet and HowNet: Unsupervised Sense Alignment and Sememe Prediction","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Machine Intelligence Institute","keywords":"Computer science; Bridging (networking); Natural language processing; Artificial intelligence; Information retrieval","score_opus":0.04240157804227087,"score_gpt":0.24002007813924375,"score_spread":0.19761850009697288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566650","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11569314,0.001637786,0.85750914,0.0006571756,0.00044609292,0.00026468356,0.0046835374,0.011890435,0.0072180163],"genre_scores_gemma":[0.4323418,0.000807394,0.5383041,0.0002604753,0.0001729496,0.00035262192,0.020456713,0.0016885714,0.005615361],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980033,0.00065044843,0.0001547934,0.000857067,0.00021989712,0.000114512666],"domain_scores_gemma":[0.99586403,0.0022854507,0.00035138152,0.0007428042,0.000591746,0.00016447541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022046356,0.001628934,0.00091900805,0.004885169,0.0011179029,0.0019955568,0.0015667791,0.0012394257,0.0033502101],"category_scores_gemma":[0.0067429007,0.00064343784,0.0012540711,0.003825699,0.00077946077,0.0057817074,0.0028779663,0.001972088,0.0027703783],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082095974,0.00044461174,0.019163135,0.0013859214,0.00055405224,0.00087546656,0.0026476681,0.02568381,0.04332013,0.028410628,0.02435927,0.85233444],"study_design_scores_gemma":[0.000092742186,0.00025616417,0.01650836,0.00039578322,0.00023758985,0.0012458367,0.0021316977,0.7895317,0.04920645,0.076628275,0.06361055,0.00015477571],"about_ca_topic_score_codex":0.0033308752,"about_ca_topic_score_gemma":0.009587037,"teacher_disagreement_score":0.004885169,"about_ca_system_score_codex":0.00063289166,"about_ca_system_score_gemma":0.0014568134,"threshold_uncertainty_score":0.011659384},"labels":[],"label_agreement":null},{"id":"W4386566659","doi":"10.18653/v1/2023.eacl-main.239","title":"DyLoRA: Parameter-Efficient Tuning of Pre-trained Models using Dynamic Search-Free Low-Rank Adaptation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Benchmark (surveying); Rank (graph theory); Sorting; Language model; Range (aeronautics); Task (project management); Machine learning; Artificial intelligence; Algorithm","score_opus":0.07253695507708174,"score_gpt":0.29376392995067,"score_spread":0.22122697487358828,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566659","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032839157,0.0013042686,0.93304074,0.00041383668,0.00022102166,0.0001984249,0.00043677236,0.028179532,0.0033662545],"genre_scores_gemma":[0.48846436,0.000545301,0.4931149,0.0013702336,0.00016829492,0.0008012952,0.003538847,0.0038458833,0.008150911],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99911386,0.00024643078,0.00006010387,0.000284135,0.00016542655,0.00012996806],"domain_scores_gemma":[0.998108,0.0009510937,0.00011568044,0.000396972,0.00031870927,0.000109517874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014956581,0.0030114027,0.0015443364,0.0008862388,0.00057191827,0.001435565,0.0034265334,0.0021700428,0.006613118],"category_scores_gemma":[0.008303456,0.0012782594,0.0015000827,0.00070537336,0.000925112,0.002310017,0.0022492989,0.005214023,0.0051155514],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003492856,0.0004437952,0.0018217566,0.0002879524,0.0003060924,0.00024271263,0.0001628141,0.63688433,0.01448427,0.003033493,0.01995481,0.32202867],"study_design_scores_gemma":[0.000027944883,0.000039187522,0.00011538577,0.000008739716,0.000013923324,0.00002154242,0.000014884718,0.996086,0.0017535007,0.0010594622,0.000848291,0.000011142844],"about_ca_topic_score_codex":0.010321283,"about_ca_topic_score_gemma":0.01884442,"teacher_disagreement_score":0.010321283,"about_ca_system_score_codex":0.0009208658,"about_ca_system_score_gemma":0.0017059125,"threshold_uncertainty_score":0.022123039},"labels":[],"label_agreement":null},{"id":"W4386566692","doi":"10.18653/v1/2023.eacl-main.150","title":"TwiRGCN: Temporally Weighted Graph Convolution for Question Answering over Temporal Knowledge Graphs","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada)","funders":"Science and Engineering Research Board","keywords":"Computer science; Knowledge graph; Question answering; Convolution (computer science); Graph; Artificial intelligence; Theoretical computer science; Artificial neural network","score_opus":0.029409841315950293,"score_gpt":0.2913399630388992,"score_spread":0.2619301217229489,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566692","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022297932,0.0028181735,0.90863043,0.0010124316,0.00041075356,0.00042761458,0.0059266156,0.05321317,0.0052629253],"genre_scores_gemma":[0.1940242,0.0011269739,0.76800424,0.0007329818,0.00019112969,0.00054592115,0.021530166,0.0018313234,0.01201301],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990827,0.00021968385,0.000060190672,0.00030017114,0.0002530784,0.00008424125],"domain_scores_gemma":[0.99881774,0.00056974415,0.000058697522,0.00027617346,0.00019815232,0.00007942843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017732853,0.0010485896,0.0013807946,0.0019658033,0.0009841331,0.0017894087,0.0024144012,0.0019705691,0.008514973],"category_scores_gemma":[0.005092343,0.00056690467,0.0013549628,0.002492832,0.00049215066,0.004724065,0.002565729,0.0018933401,0.0036328905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000848692,0.0005523956,0.0019272724,0.00054598757,0.000267252,0.00031960904,0.00042912515,0.0499823,0.009035904,0.0270618,0.120750815,0.788279],"study_design_scores_gemma":[0.00007707561,0.00008062686,0.0007617674,0.00004002987,0.000058237452,0.00012240405,0.00012466004,0.93353164,0.004231537,0.039641917,0.02130044,0.000029684647],"about_ca_topic_score_codex":0.034463212,"about_ca_topic_score_gemma":0.057514094,"teacher_disagreement_score":0.034463212,"about_ca_system_score_codex":0.0016250684,"about_ca_system_score_gemma":0.0016796547,"threshold_uncertainty_score":0.068525255},"labels":[],"label_agreement":null},{"id":"W4386566695","doi":"10.18653/v1/2023.eacl-main.206","title":"The StatCan Dialogue Dataset: Retrieving Data Tables through Conversations with Genuine Intents","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Conversation; Computer science; Task (project management); Table (database); Set (abstract data type); Data set; Information retrieval; Data science; Natural language processing; Artificial intelligence; Data mining; Linguistics","score_opus":0.11115076708320125,"score_gpt":0.3048785513780047,"score_spread":0.19372778429480342,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566695","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13871236,0.0041151256,0.0146881165,0.0023863208,0.0008992423,0.0013851823,0.79851824,0.023669755,0.015625628],"genre_scores_gemma":[0.09547059,0.0003734907,0.02188376,0.00055235473,0.00013788682,0.0007567303,0.87573355,0.0005494082,0.0045422576],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9957397,0.0017109851,0.0003120329,0.00088020606,0.00094533246,0.0004116753],"domain_scores_gemma":[0.99367213,0.002759496,0.0002957221,0.0014366612,0.0012817816,0.0005542928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023004648,0.002900877,0.0012241529,0.0033681232,0.002466442,0.0025966528,0.0032580625,0.003231575,0.008136845],"category_scores_gemma":[0.013432241,0.0005124266,0.0014530373,0.003215457,0.0010661362,0.002734634,0.0025499312,0.0026633765,0.010998648],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019108519,0.0015313124,0.021531155,0.002533542,0.00039455638,0.00097035925,0.0022071933,0.013578176,0.0056774663,0.0031078483,0.8717566,0.07480097],"study_design_scores_gemma":[0.001085792,0.000827495,0.04950504,0.0007293375,0.00031732398,0.0018171284,0.007830921,0.17490415,0.019611634,0.010107349,0.73266685,0.0005969968],"about_ca_topic_score_codex":0.10044557,"about_ca_topic_score_gemma":0.17847016,"teacher_disagreement_score":0.89955443,"about_ca_system_score_codex":0.00292976,"about_ca_system_score_gemma":0.0040363404,"threshold_uncertainty_score":0.19972181},"labels":[],"label_agreement":null},{"id":"W4386566749","doi":"10.18653/v1/2023.eacl-main.207","title":"Question Generation Using Sequence-to-Sequence Model with Semantic Role Labels","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ontario Institute of Technology; York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Sequence (biology); Sentence; Artificial intelligence; Natural language processing; Semantic role labeling; Sequence labeling; Engineering","score_opus":0.12582021380446218,"score_gpt":0.3205097142631094,"score_spread":0.19468950045864725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566749","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030364016,0.00048740514,0.9577202,0.00046634965,0.0001651944,0.00039690334,0.0010598977,0.007251125,0.0020889044],"genre_scores_gemma":[0.38906652,0.00036507516,0.5968537,0.00058671914,0.00014397435,0.0007498958,0.006347081,0.00043452947,0.0054525444],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988116,0.00046034664,0.00008677252,0.00040959936,0.00017414363,0.00005749547],"domain_scores_gemma":[0.9970355,0.0018685572,0.00014790727,0.00034814756,0.00048648744,0.00011347886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016881388,0.000987941,0.0007597129,0.0010167051,0.00035881854,0.000840227,0.0018301107,0.0015695617,0.0035833726],"category_scores_gemma":[0.0054326802,0.00036914297,0.0014636312,0.0006943118,0.00045289678,0.0020833632,0.0008500816,0.0015591877,0.0017294604],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005994958,0.00070222304,0.007259605,0.00075082405,0.00021031199,0.00087166746,0.00091062667,0.24372274,0.041555166,0.020184053,0.02406921,0.65916413],"study_design_scores_gemma":[0.000035477406,0.00008265573,0.00039530854,0.0000110933015,0.000023521658,0.000116214396,0.000041316613,0.9792117,0.007445308,0.009524438,0.0030968776,0.000016030881],"about_ca_topic_score_codex":0.004299547,"about_ca_topic_score_gemma":0.005372184,"teacher_disagreement_score":0.004299547,"about_ca_system_score_codex":0.0008075883,"about_ca_system_score_gemma":0.0011244098,"threshold_uncertainty_score":0.011987567},"labels":[],"label_agreement":null},{"id":"W4386566755","doi":"10.18653/v1/2023.eacl-main.125","title":"What happens before and after: Multi-Event Commonsense in Event Coreference Resolution","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Coreference; Computer science; Event (particle physics); Leverage (statistics); Natural language processing; Artificial intelligence; Commonsense knowledge; Commonsense reasoning; Sentence; Resolution (logic); Knowledge-based systems","score_opus":0.04628439186879934,"score_gpt":0.29076791842976474,"score_spread":0.2444835265609654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566755","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038170848,0.00038611196,0.95508116,0.0007406565,0.00008495379,0.000068040266,0.00042528697,0.0010184092,0.00402455],"genre_scores_gemma":[0.73636603,0.00028700722,0.25845936,0.0002955441,0.00013354582,0.00010590611,0.0011236466,0.00029044593,0.0029385346],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972952,0.0010290933,0.00012881683,0.0010142399,0.00037044392,0.00016211519],"domain_scores_gemma":[0.9914528,0.0062507004,0.0005100765,0.0011450375,0.00048901356,0.00015224637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037791494,0.00083662546,0.0007783326,0.0020265414,0.0017530943,0.00225292,0.0031581633,0.0023810295,0.004570787],"category_scores_gemma":[0.016764836,0.0008272282,0.0016323647,0.0019657332,0.0015274042,0.0067843036,0.0040038438,0.0028766354,0.00094239006],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012050894,0.00027379137,0.014213623,0.00066439144,0.000498921,0.003454472,0.010783709,0.22084667,0.017602088,0.3685266,0.012875086,0.3490556],"study_design_scores_gemma":[0.00003268208,0.000034682893,0.0017530416,0.00006889037,0.00012329513,0.00065315625,0.0007099426,0.76822436,0.008765472,0.20885865,0.010696571,0.00007932311],"about_ca_topic_score_codex":0.0055216243,"about_ca_topic_score_gemma":0.008404576,"teacher_disagreement_score":0.0055216243,"about_ca_system_score_codex":0.001117443,"about_ca_system_score_gemma":0.001107827,"threshold_uncertainty_score":0.019986272},"labels":[],"label_agreement":null},{"id":"W4386566844","doi":"10.18653/v1/2023.eacl-demo.10","title":"NxPlain: A Web-based Tool for Discovery of Latent Concepts","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Computational linguistics; World Wide Web; Natural language processing","score_opus":0.03701453510021448,"score_gpt":0.28585486595133724,"score_spread":0.24884033085112278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566844","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059828777,0.0015673586,0.6675485,0.0008867741,0.00026104896,0.0005921118,0.05331559,0.2627769,0.00706881],"genre_scores_gemma":[0.04804229,0.0016047468,0.7859716,0.0006057493,0.0001983378,0.0032731106,0.1355195,0.01081101,0.013973639],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984451,0.00057347544,0.00017023897,0.00033535794,0.0004223384,0.000053607935],"domain_scores_gemma":[0.9945458,0.004226456,0.0003059381,0.00046048278,0.00029247746,0.00016893251],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033011762,0.002249269,0.0012202371,0.0064767865,0.001262633,0.0027218263,0.0018889093,0.0012891305,0.03983043],"category_scores_gemma":[0.01092459,0.0010143294,0.0016485505,0.003566821,0.00047649522,0.0052010445,0.0048669116,0.0019102814,0.019138094],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013595972,0.00045570714,0.0071462207,0.0031982714,0.0006706358,0.0014760231,0.0021127088,0.00354441,0.011652009,0.024509728,0.41439945,0.5294752],"study_design_scores_gemma":[0.0010076001,0.00028104696,0.0095272735,0.0007537024,0.00032154718,0.0017582697,0.0016692829,0.22222088,0.019553736,0.13046473,0.6121236,0.0003183361],"about_ca_topic_score_codex":0.0028667222,"about_ca_topic_score_gemma":0.008160898,"teacher_disagreement_score":0.03983043,"about_ca_system_score_codex":0.00066533574,"about_ca_system_score_gemma":0.0013102045,"threshold_uncertainty_score":0.133246},"labels":[],"label_agreement":null},{"id":"W4386566845","doi":"10.18653/v1/2023.eacl-main.13","title":"Do we need Label Regularization to Fine-tune Pre-trained Language Models?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; University of Waterloo","funders":"","keywords":"Pascal (unit); Regularization (linguistics); Computer science; Natural language processing; Artificial intelligence; Theoretical computer science; Programming language","score_opus":0.03158016076292302,"score_gpt":0.2743124755449816,"score_spread":0.2427323147820586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566845","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026744636,0.004847805,0.93783146,0.011237679,0.001602955,0.00011796832,0.0009005013,0.010471123,0.0062457663],"genre_scores_gemma":[0.3593691,0.002560941,0.6033698,0.009207517,0.0019562047,0.0003972179,0.004645866,0.006197942,0.012295443],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99721706,0.0013844147,0.0000967851,0.0008039669,0.00024683215,0.0002509929],"domain_scores_gemma":[0.9891734,0.006585806,0.00036561265,0.002085294,0.0013440251,0.00044577217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00730774,0.002170084,0.0022427752,0.000820792,0.0011968793,0.0028845912,0.0035599843,0.0042845,0.00510933],"category_scores_gemma":[0.030160366,0.0012307011,0.001172033,0.001037557,0.0013299893,0.007879703,0.0018163534,0.007216375,0.0075604604],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087169016,0.0008841802,0.008285618,0.00084498385,0.0008110793,0.00020802584,0.000512924,0.0939047,0.026173566,0.021346383,0.1270219,0.7191349],"study_design_scores_gemma":[0.00022860535,0.00012328559,0.0019988862,0.0002325513,0.00020938453,0.00019029727,0.00031631486,0.86811864,0.009537104,0.09475431,0.024193263,0.00009725451],"about_ca_topic_score_codex":0.010862185,"about_ca_topic_score_gemma":0.029554438,"teacher_disagreement_score":0.010862185,"about_ca_system_score_codex":0.001039015,"about_ca_system_score_gemma":0.0022934612,"threshold_uncertainty_score":0.038647473},"labels":[],"label_agreement":null},{"id":"W4386566852","doi":"10.18653/v1/2023.eacl-demo.23","title":"CoTEVer: Chain of Thought Prompting Annotation Toolkit for Explanation Verification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Ministry of Science and ICT, South Korea; Institute for Information and Communications Technology Promotion; Yonsei University","keywords":"Computer science; Annotation; Chain (unit); Natural language processing; Human–computer interaction; Artificial intelligence","score_opus":0.05112360901474697,"score_gpt":0.28959491631619616,"score_spread":0.2384713073014492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566852","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024148654,0.00042533953,0.5463779,0.000539523,0.0005523527,0.00045069592,0.01631859,0.42482826,0.008092487],"genre_scores_gemma":[0.06623081,0.000534386,0.8005642,0.00058141246,0.00020075857,0.0013691179,0.06095804,0.051673558,0.017887715],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99813735,0.00061968155,0.00023461941,0.0004226933,0.00047451092,0.00011107735],"domain_scores_gemma":[0.9897015,0.0058186753,0.00033853433,0.0017405791,0.0020389138,0.0003618308],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036454434,0.0021345578,0.0010964073,0.00330289,0.0013794249,0.0032408717,0.0030913844,0.002167793,0.10611616],"category_scores_gemma":[0.018191544,0.0013651762,0.001916844,0.0013964004,0.0008786057,0.0054452373,0.0053445254,0.002927422,0.034538023],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010936218,0.00025642815,0.0022500479,0.0030975123,0.00021082423,0.0009791241,0.0026913143,0.0029178355,0.016615104,0.044357315,0.5989329,0.3265979],"study_design_scores_gemma":[0.0005502677,0.00015643038,0.0021832965,0.00091144367,0.00018392917,0.0009131869,0.0013241788,0.14501092,0.04337836,0.09197219,0.7131012,0.0003145917],"about_ca_topic_score_codex":0.005946417,"about_ca_topic_score_gemma":0.011416735,"teacher_disagreement_score":0.10611616,"about_ca_system_score_codex":0.0009453119,"about_ca_system_score_gemma":0.002940238,"threshold_uncertainty_score":0.35499394},"labels":[],"label_agreement":null},{"id":"W4386566893","doi":"10.18653/v1/2023.eacl-main.49","title":"Combining Parameter-efficient Modules for Task-level Generalisation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Mila - Quebec Artificial Intelligence Institute; McGill University","funders":"Institut de Valorisation des Données; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Reinforcement learning; Modular design; Benchmark (surveying); Machine learning; Artificial intelligence; Task (project management); Latent variable; Language model","score_opus":0.12780358317348253,"score_gpt":0.2946614849743965,"score_spread":0.16685790180091395,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566893","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04460303,0.00027321308,0.94854784,0.00031522123,0.000036351543,0.0001298269,0.00010875712,0.004189188,0.0017965295],"genre_scores_gemma":[0.71084565,0.00024571753,0.282375,0.00054942665,0.00008853106,0.00055640994,0.00064493425,0.0005780795,0.0041161897],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989893,0.00033444192,0.000049552258,0.00038383465,0.00012939324,0.00011346466],"domain_scores_gemma":[0.9968118,0.0013585693,0.00018543143,0.0011573471,0.00031676795,0.0001701345],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026755547,0.00196702,0.0012239133,0.00079656864,0.00046597485,0.0010952642,0.0030478393,0.002012291,0.0041177105],"category_scores_gemma":[0.010159449,0.00094491336,0.001432963,0.0007701897,0.0015998693,0.0043835645,0.003790195,0.0038594091,0.0019380214],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035699466,0.00038907726,0.004093911,0.00022690665,0.0002671791,0.00016126237,0.0003955986,0.6010508,0.020912847,0.0147860255,0.003712088,0.35364735],"study_design_scores_gemma":[0.000028276309,0.000093801886,0.00026577464,0.000011496436,0.000032146483,0.000035410336,0.000019073317,0.9801326,0.0028165306,0.0158066,0.00074508245,0.00001319513],"about_ca_topic_score_codex":0.002839242,"about_ca_topic_score_gemma":0.00544526,"teacher_disagreement_score":0.0041177105,"about_ca_system_score_codex":0.0011320432,"about_ca_system_score_gemma":0.0010656219,"threshold_uncertainty_score":0.014149845},"labels":[],"label_agreement":null},{"id":"W4386566897","doi":"10.18653/v1/2023.eacl-demo.3","title":"NLP Workbench: Efficient and Extensible Integration of State-of-the-art Text Mining Tools","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Workbench; Computer science; Natural language processing; Biomedical text mining; Extensibility; State (computer science); Artificial intelligence; Computational linguistics; Information retrieval; Programming language; Text mining; Visualization","score_opus":0.04717803324324346,"score_gpt":0.2622634332073935,"score_spread":0.21508539996415002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566897","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005897163,0.0010971903,0.53052765,0.00082957296,0.00061079743,0.00092319533,0.046589024,0.40626702,0.0072583603],"genre_scores_gemma":[0.053123757,0.0015035332,0.683837,0.0009267497,0.000331682,0.0029046661,0.21250826,0.03097976,0.013884668],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99763584,0.00052878907,0.00041218204,0.0005956135,0.00073782745,0.00008976652],"domain_scores_gemma":[0.99491453,0.0027013528,0.00021367757,0.0011178002,0.00075017684,0.00030241674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004062647,0.0029551885,0.0015121843,0.007039862,0.0013743228,0.0038051358,0.003524877,0.0013940658,0.016902167],"category_scores_gemma":[0.012311531,0.0015721516,0.001380596,0.0049234144,0.0006485722,0.0072830557,0.0053776875,0.0020637645,0.020367907],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010664616,0.00054690795,0.0027574662,0.0017216767,0.0005700695,0.0013493267,0.0013080275,0.0039260834,0.023098694,0.008870714,0.48287696,0.47190762],"study_design_scores_gemma":[0.0011017948,0.00039389322,0.0044204216,0.0005489352,0.00040518545,0.0013613824,0.0013671763,0.25815687,0.07839089,0.050466437,0.60300106,0.00038598885],"about_ca_topic_score_codex":0.0037619846,"about_ca_topic_score_gemma":0.0055343984,"teacher_disagreement_score":0.016902167,"about_ca_system_score_codex":0.0005843471,"about_ca_system_score_gemma":0.0017893575,"threshold_uncertainty_score":0.05654341},"labels":[],"label_agreement":null},{"id":"W4386566901","doi":"10.18653/v1/2023.eacl-main.19","title":"Understanding Transformer Memorization Recall Through Idioms","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"European Commission; Canadian Institute for Advanced Research","keywords":"Memorization; Computer science; Transformer; Recall; Phrase; Artificial intelligence; Natural language processing; Security token; Machine learning; Speech recognition; Cognitive psychology; Psychology","score_opus":0.237016441629294,"score_gpt":0.29891667392559856,"score_spread":0.061900232296304564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566901","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72851586,0.00030728558,0.263542,0.000601718,0.000028002749,0.00008695019,0.00031016872,0.0008057818,0.005802376],"genre_scores_gemma":[0.9672222,0.000112004076,0.03128187,0.00007750042,0.000009637383,0.00004618952,0.00031326638,0.00010858975,0.0008287262],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987508,0.00049509195,0.000083615865,0.0003953876,0.00017263068,0.00010249697],"domain_scores_gemma":[0.9858279,0.009784844,0.0012072896,0.00221383,0.0007637873,0.00020241283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035220995,0.0005458416,0.00046307564,0.0007981212,0.00037690223,0.002584094,0.0011652246,0.0008010994,0.002805076],"category_scores_gemma":[0.034754593,0.0005213706,0.00064307183,0.0004958341,0.0012881996,0.007714496,0.0018351357,0.0020780168,0.00050639524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014543707,0.00042194565,0.20873216,0.0009906277,0.00057913165,0.0011532356,0.022420878,0.11095593,0.078461386,0.1222715,0.0036772287,0.44888163],"study_design_scores_gemma":[0.00010823169,0.00061646337,0.05042077,0.00019385593,0.00032955283,0.0011134178,0.0038158281,0.6627117,0.056869436,0.21646908,0.007188105,0.00016362256],"about_ca_topic_score_codex":0.0015649445,"about_ca_topic_score_gemma":0.002036208,"teacher_disagreement_score":0.0035220995,"about_ca_system_score_codex":0.00091564865,"about_ca_system_score_gemma":0.0005238056,"threshold_uncertainty_score":0.018626869},"labels":[],"label_agreement":null},{"id":"W4386566906","doi":"10.18653/v1/2023.eacl-main.50","title":"Self-imitation Learning for Action Generation in Text-based Games","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Leverage (statistics); Computer science; Reinforcement learning; Artificial intelligence; Exploit; Pruning; Action (physics); Machine learning; Imitation; Set (abstract data type); Rank (graph theory); Mathematics","score_opus":0.0835096085731733,"score_gpt":0.3065685559575082,"score_spread":0.2230589473843349,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566906","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08456201,0.0006121595,0.90969974,0.0005603248,0.000083132974,0.00014100428,0.00010611788,0.0008845413,0.0033508933],"genre_scores_gemma":[0.9414137,0.00014484905,0.054478683,0.00021164904,0.000044448112,0.00018810725,0.00014225578,0.00011057833,0.0032657885],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992244,0.00030276921,0.000046023506,0.00020727263,0.00012579169,0.00009375078],"domain_scores_gemma":[0.9946414,0.0041327663,0.0004132554,0.00021019718,0.00035434077,0.0002480749],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015787248,0.0012552955,0.0014335233,0.0006308874,0.00042833417,0.0008973073,0.001938716,0.0013761276,0.0028358502],"category_scores_gemma":[0.009803992,0.000528028,0.0006272114,0.0003687172,0.0013101858,0.0019541292,0.0011488507,0.0018457709,0.0005350839],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012222765,0.00012533828,0.0014876103,0.000105104766,0.00005500411,0.00013298237,0.0001711806,0.94270957,0.002215457,0.012562036,0.00078366435,0.039529774],"study_design_scores_gemma":[0.000007428196,0.000023486937,0.000057585185,0.0000034315497,0.000003457013,0.000008260943,0.0000046064033,0.99603206,0.0002041755,0.0035469187,0.00010484057,0.0000037007849],"about_ca_topic_score_codex":0.0050542527,"about_ca_topic_score_gemma":0.004738271,"teacher_disagreement_score":0.0050542527,"about_ca_system_score_codex":0.00112094,"about_ca_system_score_gemma":0.0009697095,"threshold_uncertainty_score":0.010049701},"labels":[],"label_agreement":null},{"id":"W4386566909","doi":"10.18653/v1/2023.eacl-main.4","title":"Shironaam: Bengali News Headline Generation using Auxiliary Information","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Headline; Bengali; Computer science; Unavailability; Artificial intelligence; Language model; Natural language processing; Encoder; Information retrieval; Linguistics; Mathematics","score_opus":0.09463764275281533,"score_gpt":0.29184910065051817,"score_spread":0.19721145789770284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566909","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13723156,0.0054367846,0.36197844,0.0019055946,0.0032087537,0.0021725663,0.065980785,0.38314903,0.03893658],"genre_scores_gemma":[0.3296538,0.0017386408,0.39371985,0.00092513126,0.0007723586,0.0012249664,0.19241194,0.008095846,0.07145737],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931526,0.0001549394,0.000050963365,0.0002647968,0.00013760345,0.00007640809],"domain_scores_gemma":[0.998447,0.00026643943,0.000104525716,0.00052801915,0.0005339423,0.000120147495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00071024516,0.002019132,0.00075533474,0.0016835268,0.00082940934,0.0015229066,0.0017837511,0.0008722558,0.010286136],"category_scores_gemma":[0.0032569016,0.0004634775,0.0007633188,0.0013525371,0.00034048842,0.0017456325,0.0013068812,0.0011056109,0.0116508575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012908957,0.0004910988,0.0050601573,0.001444958,0.0002962767,0.0009779896,0.001205476,0.009155263,0.071421936,0.002415797,0.32372573,0.5825145],"study_design_scores_gemma":[0.00063188554,0.0007984666,0.012922169,0.00013845121,0.0004425629,0.0014896325,0.0013775363,0.31280258,0.26538336,0.0042501534,0.3994467,0.0003165496],"about_ca_topic_score_codex":0.019821646,"about_ca_topic_score_gemma":0.033297885,"teacher_disagreement_score":0.019821646,"about_ca_system_score_codex":0.0009182684,"about_ca_system_score_gemma":0.0010750946,"threshold_uncertainty_score":0.039412558},"labels":[],"label_agreement":null},{"id":"W4386566920","doi":"10.18653/v1/2023.eacl-main.42","title":"Incorporating Question Answering-Based Signals into Abstractive Summarization via Salient Span Selection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited; Advanced Research Projects Agency; Defense Advanced Research Projects Agency; U.S. Department of Defense","keywords":"Automatic summarization; Salient; Computer science; Benchmark (surveying); Question answering; Selection (genetic algorithm); Artificial intelligence; Natural language processing; Multi-document summarization; Information retrieval","score_opus":0.01698425543178204,"score_gpt":0.2674024379954873,"score_spread":0.25041818256370524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386566920","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03143598,0.0005047444,0.96122855,0.00043274154,0.00006707001,0.00018626844,0.00038156542,0.004723271,0.0010397794],"genre_scores_gemma":[0.48871776,0.00034681734,0.50232196,0.00034264507,0.00025888984,0.0004031261,0.0031142498,0.0004662062,0.0040283636],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99871504,0.0005025437,0.00008929178,0.00042465076,0.00020456292,0.00006396958],"domain_scores_gemma":[0.9957878,0.0022836875,0.00043552948,0.0005881188,0.00076492206,0.00013989523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020262105,0.0011281715,0.00085046556,0.0011742656,0.00041387044,0.0011330438,0.001634679,0.001084074,0.0023969463],"category_scores_gemma":[0.008249718,0.00039406607,0.00086120627,0.0007747414,0.0004717476,0.002809851,0.0011190568,0.0016215058,0.0014590542],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063810113,0.00045069162,0.0054386216,0.00068384164,0.00023365846,0.00027786635,0.0014422502,0.09879635,0.090494245,0.010098909,0.008175982,0.7832695],"study_design_scores_gemma":[0.00005285538,0.00058612163,0.001818504,0.00004592869,0.00013306965,0.0001337516,0.00017521762,0.9321336,0.039638534,0.015309885,0.009920527,0.00005199238],"about_ca_topic_score_codex":0.0014153813,"about_ca_topic_score_gemma":0.0025757882,"teacher_disagreement_score":0.0023969463,"about_ca_system_score_codex":0.00056365517,"about_ca_system_score_gemma":0.0008009948,"threshold_uncertainty_score":0.010715783},"labels":[],"label_agreement":null},{"id":"W4386576671","doi":"10.18653/v1/2023.vardial-1.24","title":"SIDLR: Slot and Intent Detection Models for Low-Resource Language Varieties","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Margin (machine learning); Task (project management); Dialog box; Generalization; Baseline (sea); Encoder; Natural language; Language model; Natural language understanding; Natural language processing; Resource (disambiguation); Artificial intelligence; Speech recognition; Machine learning; Engineering; World Wide Web","score_opus":0.025465777222560968,"score_gpt":0.239637146382223,"score_spread":0.21417136915966203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386576671","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09287474,0.001884136,0.8519716,0.0012296578,0.0004750918,0.00038259383,0.005342419,0.04094058,0.0048992285],"genre_scores_gemma":[0.56056345,0.0005801305,0.40802747,0.0008588745,0.00031673085,0.00073293137,0.015254552,0.0020759962,0.01158979],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985758,0.00059305015,0.00007706069,0.00043705053,0.00019470834,0.00012236086],"domain_scores_gemma":[0.99635255,0.002358221,0.00012572319,0.0005848808,0.0004285428,0.00014999158],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035872208,0.0022545562,0.0011153638,0.0012082418,0.0007369638,0.001730144,0.0030182807,0.0017960476,0.00648508],"category_scores_gemma":[0.0077165104,0.0007944302,0.0018131532,0.00073339953,0.0006207648,0.003875262,0.0027155005,0.0042708437,0.005119242],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017441981,0.0008667061,0.008283926,0.0007762575,0.0003954272,0.00052880053,0.001323209,0.15363298,0.024563044,0.015780209,0.07165858,0.72044665],"study_design_scores_gemma":[0.00006194302,0.00014981416,0.0009237686,0.000025921472,0.000044944205,0.00017091341,0.00013285562,0.9744505,0.0058559515,0.011581102,0.0065454296,0.000056780074],"about_ca_topic_score_codex":0.007443355,"about_ca_topic_score_gemma":0.010883381,"teacher_disagreement_score":0.007443355,"about_ca_system_score_codex":0.0010274614,"about_ca_system_score_gemma":0.0012450977,"threshold_uncertainty_score":0.02169472},"labels":[],"label_agreement":null},{"id":"W4386576803","doi":"10.18653/v1/2023.findings-eacl.106","title":"More Robust Schema-Guided Dialogue State Tracking via Tree-Based Paraphrase Ranking","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Toyota Motor Europe; Atomic Energy of Canada Limited; Department of Science and Technology, Ministry of Science and Technology, India","keywords":"Computer science; Schema (genetic algorithms); Artificial intelligence; Natural language processing; Paraphrase; Scalability; Natural language understanding; Robustness (evolution); Machine learning; Information retrieval; Natural language; Database","score_opus":0.06953959777686648,"score_gpt":0.2812464530772721,"score_spread":0.21170685530040562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386576803","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12671,0.0014663489,0.8111115,0.0007358071,0.00034132847,0.0004521117,0.003656914,0.046878744,0.008647252],"genre_scores_gemma":[0.54420185,0.000301923,0.43459132,0.00038566248,0.00009625608,0.00033703764,0.013270735,0.0014630421,0.0053521623],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9961436,0.0019827394,0.00019500304,0.0010453954,0.00045921683,0.00017400923],"domain_scores_gemma":[0.9938147,0.0029862698,0.00025791358,0.0018378515,0.00088884553,0.0002144092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039269193,0.001342551,0.0013762416,0.0014140476,0.0005887235,0.0026641656,0.0024740442,0.0015596465,0.0044018594],"category_scores_gemma":[0.0146010285,0.0004526812,0.0009127668,0.0013068041,0.00045859456,0.0042336956,0.0022191035,0.0024185565,0.004459696],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008205806,0.0010427897,0.0057436167,0.00074763893,0.0003160383,0.0001826945,0.0010556434,0.11324639,0.044688288,0.007011954,0.037187874,0.7879566],"study_design_scores_gemma":[0.0001187071,0.00025306144,0.001031421,0.00003244925,0.000054827793,0.0001139274,0.00024245294,0.96449417,0.019072676,0.007139248,0.0073978435,0.00004918755],"about_ca_topic_score_codex":0.0051791035,"about_ca_topic_score_gemma":0.0077984473,"teacher_disagreement_score":0.0051791035,"about_ca_system_score_codex":0.0007163983,"about_ca_system_score_gemma":0.0015834896,"threshold_uncertainty_score":0.020767748},"labels":[],"label_agreement":null},{"id":"W4386603805","doi":"10.3233/faia230232","title":"Knowledge of Language and Natural Language Processing","year":2023,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Natural language processing; Universal Networking Language; Language identification; Question answering; Artificial intelligence; Natural language programming; Natural language; Anaphora (linguistics); Covert; Dependency (UML); Sentence; Object language; Generative grammar; Linguistics; Comprehension approach","score_opus":0.036793372426872716,"score_gpt":0.3008611652503967,"score_spread":0.264067792823524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386603805","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0052128015,0.1514756,0.16974524,0.029544996,0.0015728549,0.0000616133,0.00041196402,0.0009526201,0.6410224],"genre_scores_gemma":[0.255984,0.18522717,0.14228535,0.008398729,0.009348414,0.00036251897,0.0016856738,0.0010767288,0.3956315],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9995741,0.00012921516,0.000017386465,0.00009773299,0.00015756881,0.000024008565],"domain_scores_gemma":[0.9983045,0.0014057204,0.000036135123,0.00014975222,0.00007115357,0.000032711527],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008343514,0.000583037,0.0004568523,0.0013453439,0.00060962647,0.0053603514,0.0008258274,0.0013401099,0.011068705],"category_scores_gemma":[0.0025931578,0.00039174253,0.00037222638,0.0016532541,0.0048702727,0.00871704,0.0013080108,0.0026801736,0.0044241627],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000075765042,0.0000117384125,0.000112300564,0.00024930228,0.000008103804,0.000035651403,0.0005867237,0.0015627324,0.00037920385,0.86360425,0.032976337,0.10046607],"study_design_scores_gemma":[0.0000020954853,0.0000041417434,0.00019544538,0.000167016,0.00000412632,0.00008072585,0.00010484428,0.0027428067,0.00018345323,0.8055929,0.19091372,0.000008758493],"about_ca_topic_score_codex":0.0015592107,"about_ca_topic_score_gemma":0.001452692,"teacher_disagreement_score":0.011068705,"about_ca_system_score_codex":0.0019735666,"about_ca_system_score_gemma":0.0010804356,"threshold_uncertainty_score":0.03702849},"labels":[],"label_agreement":null},{"id":"W4386712343","doi":"10.1007/978-3-031-40953-0_35","title":"Can Large Language Models Assist in Hazard Analysis?","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Hazard; Context (archaeology); Session (web analytics); Hazard analysis; Variety (cybernetics); Risk analysis (engineering); Computer security; Artificial intelligence; World Wide Web; Reliability engineering; Engineering; Ecology; Medicine","score_opus":0.02489166734725026,"score_gpt":0.2631933300556764,"score_spread":0.23830166270842615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386712343","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013475918,0.00054474716,0.97240746,0.0020854848,0.00023633362,0.00005225083,0.0011356473,0.003689128,0.006372981],"genre_scores_gemma":[0.55638415,0.0012744816,0.41641667,0.0009307772,0.00056994375,0.00033434914,0.0029540204,0.0021289696,0.01900667],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998353,0.001005583,0.0000539237,0.00019963505,0.00024607457,0.00014181188],"domain_scores_gemma":[0.9819381,0.015396121,0.00061628193,0.0010720353,0.0007172179,0.00026021307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004276797,0.0012264753,0.0011562317,0.0011958111,0.00072613277,0.0030067335,0.0022318105,0.0019783012,0.017668266],"category_scores_gemma":[0.035180308,0.0010767654,0.0018305376,0.0014258812,0.0007738427,0.007144297,0.0015671871,0.0032622819,0.005993525],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010083523,0.00036636402,0.008836082,0.00058224285,0.00042537347,0.00036594932,0.0007042771,0.35971457,0.005108636,0.23227546,0.042786427,0.34782627],"study_design_scores_gemma":[0.000057205307,0.00004610383,0.0006086564,0.000043309337,0.00007268356,0.000063966596,0.00012054237,0.7620501,0.0011347869,0.2277716,0.007998166,0.00003282159],"about_ca_topic_score_codex":0.006231613,"about_ca_topic_score_gemma":0.008914337,"teacher_disagreement_score":0.017668266,"about_ca_system_score_codex":0.00086813193,"about_ca_system_score_gemma":0.0017842922,"threshold_uncertainty_score":0.05910629},"labels":[],"label_agreement":null},{"id":"W4386724574","doi":"10.31234/osf.io/x2f4a","title":"Shadows of wisdom: Classifying meta-cognitive and morally-grounded narrative content via Large Language Models","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Waterloo","funders":"","keywords":"Narrative; Humility; Psychology; Categorization; Classifier (UML); Workflow; Social psychology; Computer science; Artificial intelligence; Linguistics; Political science","score_opus":0.23135779496857703,"score_gpt":0.32765282387620026,"score_spread":0.09629502890762323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386724574","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6674873,0.00082472415,0.3160304,0.0013945932,0.00010118,0.0007049253,0.004486202,0.003507773,0.005462847],"genre_scores_gemma":[0.8524138,0.00014263864,0.14191274,0.000113883485,0.000028067803,0.00035037592,0.0037844717,0.00019003883,0.0010639894],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99717796,0.0017845707,0.0001473521,0.00046352888,0.00031570275,0.00011078374],"domain_scores_gemma":[0.9706654,0.024265457,0.0013183156,0.0017319605,0.0015483724,0.00047047588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007049048,0.00096323167,0.0003812576,0.0027759098,0.0009126024,0.0030000668,0.0013988308,0.0008632639,0.0016767875],"category_scores_gemma":[0.03456535,0.0003600329,0.0011469753,0.0010823807,0.0010158675,0.0035757737,0.0023042161,0.0018212992,0.00079440937],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001768333,0.00070606376,0.17005031,0.0018493526,0.0005818772,0.0009045035,0.06834864,0.07482251,0.022278894,0.028514031,0.015234354,0.61494106],"study_design_scores_gemma":[0.000075303265,0.00018824264,0.03210264,0.00030286025,0.00018053618,0.0002739839,0.01025435,0.8895992,0.011959379,0.0395393,0.015372992,0.00015122331],"about_ca_topic_score_codex":0.014106688,"about_ca_topic_score_gemma":0.02386384,"teacher_disagreement_score":0.014106688,"about_ca_system_score_codex":0.0022156632,"about_ca_system_score_gemma":0.0020152559,"threshold_uncertainty_score":0.037279427},"labels":[],"label_agreement":null},{"id":"W4386763311","doi":"10.36227/techrxiv.24143706.v1","title":"GPT-4 as a Twitter Data Annotator: Unraveling Its Performance on a Stance Classification Task","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University; York University","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Task (project management); Computer science; Natural language processing; Artificial intelligence; Engineering","score_opus":0.2535944618526218,"score_gpt":0.34867202131662733,"score_spread":0.09507755946400553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386763311","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.71503156,0.0065947687,0.18846697,0.010510307,0.0030917414,0.0017047357,0.01797102,0.030277537,0.026351264],"genre_scores_gemma":[0.7734392,0.0007453525,0.18623401,0.0017761136,0.0005407449,0.0009233287,0.025145173,0.0012511313,0.009945039],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9909768,0.005479077,0.00036168788,0.0018624073,0.0008820391,0.0004379743],"domain_scores_gemma":[0.9765513,0.016271684,0.00077697966,0.002446893,0.0029557243,0.0009974557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017392403,0.002594754,0.0014496055,0.0031184945,0.0020912536,0.002998215,0.002872543,0.0044524507,0.0035136596],"category_scores_gemma":[0.03182294,0.0006766766,0.0014247368,0.0022045285,0.0011290901,0.004922445,0.0040812767,0.003962383,0.005130668],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0069063827,0.00217453,0.08867049,0.0021231547,0.0016081521,0.0008415778,0.0033869979,0.092878975,0.029693263,0.0037734457,0.11380818,0.6541349],"study_design_scores_gemma":[0.00023109013,0.00064815214,0.012465166,0.00014999523,0.00021195892,0.00020481457,0.0012823871,0.95746267,0.010270281,0.004282848,0.012641181,0.00014952848],"about_ca_topic_score_codex":0.029215425,"about_ca_topic_score_gemma":0.04278245,"teacher_disagreement_score":0.029215425,"about_ca_system_score_codex":0.002408095,"about_ca_system_score_gemma":0.0028089455,"threshold_uncertainty_score":0.091980875},"labels":[],"label_agreement":null},{"id":"W4386783354","doi":"10.1016/j.eswa.2023.121542","title":"Nbias: A natural language processing framework for BIAS identification in text","year":2023,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":67,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Vector Institute","funders":"Vector Institute; Government of Ontario; Canadian Institute for Advanced Research","keywords":"Computer science; Security token; Transformer; Identification (biology); Artificial intelligence; Variety (cybernetics); Natural language processing; Data science; Machine learning; Data mining; Computer security","score_opus":0.03851822480585686,"score_gpt":0.3238614681034604,"score_spread":0.2853432432976035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386783354","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012682867,0.00018711065,0.98940825,0.000179105,0.000063972,0.00012918215,0.0010617765,0.007227214,0.0004751197],"genre_scores_gemma":[0.052540176,0.00032451417,0.9376739,0.00028699604,0.00028762154,0.00062146113,0.0043119066,0.0014285797,0.0025248763],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9952187,0.0019578848,0.000460486,0.0010000818,0.0011376716,0.0002252405],"domain_scores_gemma":[0.98767567,0.007404427,0.0008877769,0.0013309201,0.0023044755,0.0003967218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007843419,0.0015818652,0.0015617075,0.0044527785,0.0016470882,0.0034572212,0.0026329746,0.0017220475,0.009321353],"category_scores_gemma":[0.02289441,0.00097640615,0.0020918597,0.002510108,0.0010762699,0.005274282,0.003410903,0.0031405166,0.0064195776],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001029109,0.00036513276,0.0057419157,0.0018076795,0.0005018627,0.0004819155,0.0019752823,0.022815723,0.028061023,0.14017734,0.06128423,0.7357588],"study_design_scores_gemma":[0.00013873525,0.00014456344,0.0019379497,0.0002445545,0.00022219084,0.00040698858,0.00038874472,0.7047128,0.018394394,0.20419382,0.06909856,0.00011670453],"about_ca_topic_score_codex":0.0052539892,"about_ca_topic_score_gemma":0.007923733,"teacher_disagreement_score":0.009321353,"about_ca_system_score_codex":0.0012710437,"about_ca_system_score_gemma":0.003058372,"threshold_uncertainty_score":0.04148048},"labels":[],"label_agreement":null},{"id":"W4386805215","doi":"10.1016/j.jbi.2023.104486","title":"A self-supervised language model selection strategy for biomedical question answering","year":2023,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Waterloo","funders":"","keywords":"Computer science; Leverage (statistics); Artificial intelligence; Classifier (UML); Machine learning; Language model; Question answering; Retraining; Domain (mathematical analysis); Task (project management); Natural language processing","score_opus":0.029768193247086988,"score_gpt":0.30776333495048175,"score_spread":0.2779951417033948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386805215","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017544057,0.0005233087,0.9751521,0.00045706323,0.00011508795,0.00019130675,0.0006012279,0.004698533,0.0007173549],"genre_scores_gemma":[0.3384112,0.00041231562,0.6433012,0.0010108773,0.00061484723,0.00080565765,0.008536793,0.0008686908,0.006038349],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966438,0.0016260756,0.00027631,0.0007274629,0.0005220174,0.00020430791],"domain_scores_gemma":[0.9936453,0.0041395323,0.0001749979,0.00055755186,0.0012690136,0.00021366912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004926242,0.0013303149,0.0018430776,0.0027037966,0.0011707006,0.0013942098,0.002834563,0.002274892,0.0030318853],"category_scores_gemma":[0.008215294,0.00071629824,0.002028528,0.0016143973,0.0005523678,0.0021833002,0.0021300875,0.0024627398,0.0028778375],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010100716,0.0010459591,0.0038168523,0.00039593503,0.00069834694,0.0004414515,0.0004806351,0.05287415,0.036496203,0.007269218,0.035114463,0.8603568],"study_design_scores_gemma":[0.000059559563,0.00010172939,0.0005821037,0.000013951215,0.00010116777,0.0001422382,0.000059410933,0.9827771,0.0073115528,0.006710719,0.002114009,0.000026391162],"about_ca_topic_score_codex":0.003374189,"about_ca_topic_score_gemma":0.00730303,"teacher_disagreement_score":0.004926242,"about_ca_system_score_codex":0.00063974265,"about_ca_system_score_gemma":0.0018860158,"threshold_uncertainty_score":0.026052833},"labels":[],"label_agreement":null},{"id":"W4386858581","doi":"10.1109/icsp58490.2023.10248908","title":"Research on Zero-Shot Stance Detection: Using Dual-Module adversarial training","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Optech (Canada)","funders":"Natural Science Foundation of Jiangxi Province; National Natural Science Foundation of China","keywords":"Adversarial system; Computer science; Shot (pellet); Dual (grammatical number); Invariant (physics); Zero (linguistics); Artificial intelligence; One shot; Training (meteorology); Training set; Baseline (sea); Machine learning; Pattern recognition (psychology); Mathematics; Engineering; Linguistics","score_opus":0.3411924000323471,"score_gpt":0.40752976174453115,"score_spread":0.06633736171218407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386858581","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15858608,0.004901645,0.8251358,0.0009873109,0.00044642753,0.0002065206,0.0004295537,0.0024403692,0.006866331],"genre_scores_gemma":[0.83951414,0.0011496708,0.14602037,0.0006115156,0.00037403274,0.00011778838,0.0017056902,0.0002741819,0.010232653],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991148,0.0003366286,0.00004154892,0.00028338836,0.00013174485,0.00009186614],"domain_scores_gemma":[0.99684674,0.0019118356,0.00023322985,0.0004993729,0.00033666167,0.00017211156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024002644,0.0014691141,0.0011604301,0.0010146464,0.00041296834,0.00091736246,0.0018447347,0.0013927593,0.0015148455],"category_scores_gemma":[0.005288127,0.0004604031,0.0008214442,0.00072071236,0.0010848726,0.0025393446,0.0013282468,0.0023025977,0.0010711086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012568012,0.0007477284,0.012177862,0.0004742121,0.00035356084,0.00036880217,0.0005001613,0.21034916,0.03841599,0.017765217,0.015517163,0.7020734],"study_design_scores_gemma":[0.000022796421,0.00020519028,0.000956634,0.0000308686,0.000032076598,0.00013691244,0.000043414162,0.9849235,0.0054277494,0.006047705,0.0021574595,0.000015768528],"about_ca_topic_score_codex":0.0014363376,"about_ca_topic_score_gemma":0.0019269264,"teacher_disagreement_score":0.0024002644,"about_ca_system_score_codex":0.0007235539,"about_ca_system_score_gemma":0.00065825315,"threshold_uncertainty_score":0.012693942},"labels":[],"label_agreement":null},{"id":"W4386870537","doi":"10.1007/978-981-99-6207-5_21","title":"Adder Encoder for Pre-trained Language Model","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Adder; Benchmark (surveying); Transformer; Encoder; Efficient energy use; Energy consumption; Language model; Computer engineering; Artificial intelligence; Electrical engineering","score_opus":0.029698172063888126,"score_gpt":0.283785682589629,"score_spread":0.25408751052574086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386870537","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017958717,0.0026371656,0.84191436,0.0014568041,0.002880151,0.0003890793,0.018973373,0.07118147,0.042608906],"genre_scores_gemma":[0.22891921,0.002476334,0.5589752,0.0014700295,0.0006285766,0.00075206667,0.038166534,0.0041500703,0.16446199],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99980503,0.000018903633,0.000014223981,0.0000686783,0.00005380914,0.000039460378],"domain_scores_gemma":[0.99968183,0.00008186521,0.000009311674,0.00007947136,0.00012824174,0.000019281966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00023321345,0.0010790759,0.00063280447,0.0005934867,0.0004609272,0.0009885022,0.0011314626,0.0008548252,0.05866915],"category_scores_gemma":[0.00091307337,0.0005132946,0.00063958904,0.0005487752,0.00020228347,0.0014516914,0.0009832582,0.0016040305,0.034218118],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045830067,0.0001276703,0.0004466206,0.00057387265,0.00007728817,0.0003492854,0.00007632336,0.013741702,0.064317934,0.024787651,0.16553435,0.729509],"study_design_scores_gemma":[0.00016474504,0.0002887159,0.0012456254,0.00022916755,0.00023963016,0.00096067536,0.00011646531,0.5087266,0.2164261,0.04473505,0.22674885,0.00011837183],"about_ca_topic_score_codex":0.007206737,"about_ca_topic_score_gemma":0.0142177865,"teacher_disagreement_score":0.05866915,"about_ca_system_score_codex":0.00069454144,"about_ca_system_score_gemma":0.0017540891,"threshold_uncertainty_score":0.19626784},"labels":[],"label_agreement":null},{"id":"W4386932912","doi":"10.1007/978-3-031-44192-9_4","title":"DaCon: Multi-Domain Text Classification Using Domain Adversarial Contrastive Learning","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Domain (mathematical analysis); Artificial intelligence; Embedding; Feature (linguistics); Encoder; Adversarial system; Class (philosophy); Pattern recognition (psychology); Machine learning; Training set; Natural language processing; Mathematics","score_opus":0.048631711654022206,"score_gpt":0.27815126678308966,"score_spread":0.22951955512906747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386932912","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012332564,0.0010007403,0.96415275,0.00043288563,0.00049465266,0.0002438728,0.0011051086,0.015810918,0.0044264942],"genre_scores_gemma":[0.22379513,0.00079763756,0.72889376,0.0010082002,0.00046445057,0.0005952544,0.008525524,0.0014601331,0.03445986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99898237,0.00023540099,0.000038275964,0.0003263482,0.00030713997,0.00011045404],"domain_scores_gemma":[0.9985234,0.0007437139,0.00006420224,0.00030333715,0.0002761013,0.000089234476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001657451,0.0014604387,0.0013177363,0.0017704539,0.0007687891,0.0014809881,0.002391153,0.0018886512,0.008495005],"category_scores_gemma":[0.0031858783,0.000502072,0.0011077669,0.0013919445,0.00065776147,0.0024046206,0.0030013637,0.003028866,0.005669989],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051253784,0.00040516126,0.00074720243,0.0002022622,0.00015945495,0.00020224079,0.000086776214,0.073327884,0.017492067,0.011512053,0.06933266,0.82601964],"study_design_scores_gemma":[0.00002548983,0.000059078055,0.00019115003,0.000012152223,0.000016116668,0.00007156842,0.0000211527,0.9768942,0.0075339973,0.008571004,0.006589446,0.00001462647],"about_ca_topic_score_codex":0.002976201,"about_ca_topic_score_gemma":0.005311341,"teacher_disagreement_score":0.008495005,"about_ca_system_score_codex":0.0008705184,"about_ca_system_score_gemma":0.0010693372,"threshold_uncertainty_score":0.0284186},"labels":[],"label_agreement":null},{"id":"W4386942346","doi":"10.1145/3624918.3625336","title":"Retrieving Supporting Evidence for Generative Question Answering","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":59,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Question answering; Computer science; Statement (logic); Hallucinating; Information retrieval; Pipeline (software); Natural language processing; Questions and answers; Artificial intelligence; Domain (mathematical analysis); Linguistics; Programming language","score_opus":0.22301359317991584,"score_gpt":0.40473926011662215,"score_spread":0.1817256669367063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386942346","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21724637,0.0017138877,0.7363096,0.005599395,0.0003207618,0.000778504,0.004372811,0.021405218,0.012253504],"genre_scores_gemma":[0.69141334,0.0002955322,0.2943428,0.00084761233,0.00023718952,0.00031006557,0.009718957,0.00080082857,0.002033711],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9904957,0.005202154,0.0005786091,0.0013048706,0.002098337,0.0003204355],"domain_scores_gemma":[0.9342041,0.05226269,0.0021776771,0.0068686297,0.0038386814,0.000648335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008665595,0.0014179439,0.00095991284,0.0037039837,0.0008232932,0.0031512969,0.002618443,0.002792542,0.011028976],"category_scores_gemma":[0.08886425,0.00076430076,0.0015900325,0.0011687941,0.0020891798,0.0060360394,0.0047006784,0.0023181017,0.0036712373],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001456225,0.00091324555,0.033047486,0.0044956617,0.00050167256,0.0039690654,0.0073949927,0.056349743,0.07395045,0.10361535,0.04649729,0.6678088],"study_design_scores_gemma":[0.00034265188,0.000360104,0.007668428,0.0005729027,0.00025464076,0.0014872974,0.0016393084,0.7477907,0.05013368,0.14924364,0.04036129,0.00014542283],"about_ca_topic_score_codex":0.0016827727,"about_ca_topic_score_gemma":0.0036265238,"teacher_disagreement_score":0.011028976,"about_ca_system_score_codex":0.001174268,"about_ca_system_score_gemma":0.0016532446,"threshold_uncertainty_score":0.04582858},"labels":[],"label_agreement":null},{"id":"W4387031631","doi":"10.48550/arxiv.2309.12886","title":"Implementing Automated Data Validation for Canadian Political Datasets","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Consistency (knowledge bases); Politics; Workflow; Government (linguistics); Work (physics); Sample (material); Computer science; Data science; Data mining; Political science; Database; Engineering; Artificial intelligence","score_opus":0.2610642722276903,"score_gpt":0.27174824669229786,"score_spread":0.01068397446460756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387031631","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23948298,0.0020895915,0.47104317,0.010027356,0.001240777,0.0048264926,0.15846464,0.08994463,0.022880368],"genre_scores_gemma":[0.27778795,0.0003328888,0.47782844,0.0016886572,0.00012885281,0.0028033026,0.23175023,0.004071768,0.0036080077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9474557,0.022344982,0.004569988,0.00761595,0.015266343,0.002747017],"domain_scores_gemma":[0.8079012,0.07924219,0.005757154,0.048745792,0.054822206,0.0035314148],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07333781,0.0017616751,0.00123668,0.00712888,0.0054887407,0.005270131,0.006005579,0.0016209494,0.0030199399],"category_scores_gemma":[0.20309655,0.0013495055,0.0021439705,0.00888634,0.002381998,0.003967556,0.0061470754,0.0040501733,0.0022959427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013293219,0.0011446675,0.20495753,0.0013591519,0.0012837328,0.0005760717,0.007570298,0.07353732,0.01675288,0.03072514,0.31480536,0.3459585],"study_design_scores_gemma":[0.0008174518,0.00020976948,0.09863637,0.0006814789,0.0002470965,0.0003090371,0.004122533,0.51486844,0.048642676,0.03401706,0.29698297,0.00046516716],"about_ca_topic_score_codex":0.5391106,"about_ca_topic_score_gemma":0.56546634,"teacher_disagreement_score":0.9266622,"about_ca_system_score_codex":0.011397747,"about_ca_system_score_gemma":0.02790179,"threshold_uncertainty_score":0.92720735},"labels":[],"label_agreement":null},{"id":"W4387107464","doi":"10.1016/j.eswa.2023.121720","title":"Parallel inference for cross-collection latent generalized Dirichlet allocation model and applications","year":2023,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Latent Dirichlet allocation; Computer science; Inference; Topic model; Hierarchical Dirichlet process; Prior probability; Dirichlet distribution; Scalability; Machine learning; Data mining; Artificial intelligence; Mathematics; Bayesian probability; Database","score_opus":0.04492236747625851,"score_gpt":0.31832813756510325,"score_spread":0.27340577008884476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387107464","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005634527,0.00045876318,0.99129206,0.0003229021,0.000111187255,0.00006509414,0.00031108028,0.0012035454,0.0006008789],"genre_scores_gemma":[0.19945017,0.00092663703,0.78454804,0.00054732856,0.00063379126,0.0007696313,0.004025059,0.001092028,0.008007443],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941817,0.002994392,0.00032555984,0.0014928603,0.00065939105,0.0003460401],"domain_scores_gemma":[0.987872,0.0077298926,0.00032042293,0.002574112,0.0011894181,0.00031412859],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009639892,0.0016562681,0.0037387311,0.0023003325,0.0019905772,0.0030593346,0.0052486164,0.002193905,0.008430622],"category_scores_gemma":[0.026390193,0.0022027034,0.0032987548,0.004060421,0.0015805343,0.0055399886,0.0043800673,0.004794232,0.0030999475],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010296437,0.0005692597,0.0035297647,0.000488135,0.00080907653,0.00027363698,0.00058920705,0.40024447,0.0035259863,0.10689805,0.018335927,0.46370685],"study_design_scores_gemma":[0.000059035912,0.000018790934,0.00029028332,0.000012486211,0.000056609657,0.000035182635,0.000031132775,0.92689055,0.0006144387,0.07042543,0.0015462484,0.000019911942],"about_ca_topic_score_codex":0.022462837,"about_ca_topic_score_gemma":0.037946925,"teacher_disagreement_score":0.022462837,"about_ca_system_score_codex":0.0023056045,"about_ca_system_score_gemma":0.0043443493,"threshold_uncertainty_score":0.050981224},"labels":[],"label_agreement":null},{"id":"W4387116260","doi":"10.1186/s40634-023-00662-4","title":"Accelerated evidence synthesis in orthopaedics—the roles of natural language processing, expert annotation and large language models","year":2023,"lang":"en","type":"article","venue":"Journal of Experimental Orthopaedics","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"Göteborgs Universitet","keywords":"Annotation; Natural language processing; Computer science; Natural language; Data science; Linguistics; Artificial intelligence; Philosophy","score_opus":0.03212934952285237,"score_gpt":0.32332320411389515,"score_spread":0.2911938545910428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387116260","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021505998,0.026730618,0.8962682,0.03618824,0.002035658,0.0013747141,0.005024128,0.0038465715,0.007025837],"genre_scores_gemma":[0.22204423,0.007223412,0.759297,0.0022707786,0.0013323317,0.0012398906,0.003145846,0.0012635827,0.002182838],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.90035754,0.080513835,0.0054826178,0.0053577614,0.007817745,0.00047054217],"domain_scores_gemma":[0.31177843,0.62863344,0.011807138,0.024085531,0.02119215,0.002503371],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13774917,0.0014906889,0.002799762,0.008316714,0.0015901088,0.011719882,0.003484219,0.003475546,0.0136670135],"category_scores_gemma":[0.5042319,0.0022720583,0.0030410618,0.0045687044,0.0028152715,0.0109495185,0.0056714755,0.005938631,0.0027956883],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00202383,0.00040425773,0.008857798,0.012998649,0.0029455048,0.0006104232,0.003677386,0.029549694,0.0049370704,0.068991676,0.04844787,0.81655586],"study_design_scores_gemma":[0.0008992659,0.00042040786,0.0073793544,0.00401886,0.0014899949,0.0005839789,0.0015676158,0.20256191,0.007156605,0.69811594,0.075379945,0.00042607865],"about_ca_topic_score_codex":0.0037066787,"about_ca_topic_score_gemma":0.0078409435,"teacher_disagreement_score":0.8622508,"about_ca_system_score_codex":0.002660853,"about_ca_system_score_gemma":0.014633556,"threshold_uncertainty_score":0.72849596},"labels":[],"label_agreement":null},{"id":"W4387123803","doi":"10.1109/rew57809.2023.00022","title":"Automatic Domain-Specific Corpora Generation from Wikipedia - A Replication Study","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Baseline (sea); Crawling; Encoder; Artificial intelligence; Natural language processing; Workflow; Replication (statistics); Replicate; Domain (mathematical analysis); Realm; World Wide Web; Information retrieval; Database","score_opus":0.0838222573724981,"score_gpt":0.2817979867145097,"score_spread":0.19797572934201157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387123803","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80219895,0.00183785,0.13608989,0.0014300296,0.0014777782,0.009390168,0.016361495,0.011230347,0.019983567],"genre_scores_gemma":[0.68386257,0.00057404814,0.2608918,0.00116102,0.000265053,0.010789826,0.033167724,0.0029183251,0.006369735],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9786309,0.012615199,0.0023254005,0.0036880854,0.0022802409,0.0004601267],"domain_scores_gemma":[0.88028145,0.043373264,0.0026051805,0.054430116,0.018257326,0.0010527316],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.023987236,0.0012390394,0.0009303547,0.0022404322,0.0014865906,0.002205538,0.0024207213,0.0015647276,0.00292771],"category_scores_gemma":[0.09555209,0.00076199684,0.0013968959,0.002027474,0.0015070554,0.0036266996,0.0030824493,0.0021165563,0.0024300735],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0065961042,0.013653777,0.07432155,0.008314174,0.0020831048,0.004061973,0.030080996,0.025483657,0.09261643,0.010912174,0.08851674,0.6433593],"study_design_scores_gemma":[0.0065035922,0.0147170145,0.11484584,0.0017181663,0.002958045,0.0063448995,0.022314448,0.11307597,0.22670382,0.025337452,0.46436134,0.0011193049],"about_ca_topic_score_codex":0.00449543,"about_ca_topic_score_gemma":0.00522045,"teacher_disagreement_score":0.97601277,"about_ca_system_score_codex":0.0009143762,"about_ca_system_score_gemma":0.0017504601,"threshold_uncertainty_score":0.12685812},"labels":[],"label_agreement":null},{"id":"W4387171785","doi":"10.3233/faia230385","title":"Diversified Prior Knowledge Enhanced General Language Model for Biomedical Information Retrieval","year":2023,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Ranking (information retrieval); Language model; Domain (mathematical analysis); Information retrieval; Query expansion; Domain knowledge; Artificial intelligence; Natural language processing","score_opus":0.052357152675397735,"score_gpt":0.29650384609911645,"score_spread":0.24414669342371872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387171785","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009953502,0.0060204156,0.97855955,0.0006820148,0.00012734637,0.00006469709,0.0005388111,0.0015513647,0.0025023052],"genre_scores_gemma":[0.34245276,0.010032131,0.6158856,0.0013721782,0.0007540361,0.00046590873,0.005699466,0.00054766575,0.022790156],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992963,0.00026901168,0.00004805341,0.00014905252,0.00019453243,0.000043013795],"domain_scores_gemma":[0.99898595,0.0006056982,0.000074020856,0.00013433058,0.00016741898,0.000032548014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012859593,0.00088205404,0.0009279546,0.0015351811,0.00024926008,0.0011192865,0.0012006949,0.0009605167,0.0024568313],"category_scores_gemma":[0.003282562,0.0003269154,0.0012473232,0.0019572284,0.00049375347,0.0021945303,0.0009740206,0.0015982784,0.002607135],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002541375,0.00018011288,0.0012679144,0.0006889594,0.00022202719,0.00039538526,0.00024453734,0.1383425,0.025920097,0.02676131,0.031118654,0.7746043],"study_design_scores_gemma":[0.000022753931,0.00011870224,0.0007926149,0.000044588924,0.00008834199,0.0003994043,0.000046588975,0.9407591,0.0052766968,0.039225657,0.013165078,0.000060507948],"about_ca_topic_score_codex":0.0031032145,"about_ca_topic_score_gemma":0.0039371084,"teacher_disagreement_score":0.0031032145,"about_ca_system_score_codex":0.00071450963,"about_ca_system_score_gemma":0.0009888799,"threshold_uncertainty_score":0.008218944},"labels":[],"label_agreement":null},{"id":"W4387172153","doi":"10.3233/faia230266","title":"From Intermediate Representations to Explanations: Exploring Hierarchical Structures in NLP","year":2023,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates; Alberta Machine Intelligence Institute","keywords":"Computer science; Artificial intelligence; Natural language processing; Abstraction; Word (group theory); Identification (biology); Interpretation (philosophy); Feature (linguistics); Class (philosophy); Linguistics","score_opus":0.11743316923527714,"score_gpt":0.31661788095002197,"score_spread":0.19918471171474483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387172153","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03536804,0.0024077941,0.94740736,0.0015622383,0.00005363234,0.00007858792,0.0011221822,0.0020538918,0.009946308],"genre_scores_gemma":[0.37227553,0.0023373023,0.6146883,0.00027408317,0.0000865686,0.00020492087,0.0033772148,0.00040553595,0.0063505564],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99969244,0.00015542892,0.000013069212,0.0000709433,0.00005009654,0.00001806276],"domain_scores_gemma":[0.99798834,0.0016689139,0.00008622786,0.00015769494,0.00006121069,0.00003747899],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00067856547,0.0008234167,0.00042349097,0.00092793116,0.00044091127,0.0023240976,0.0012052171,0.0008788736,0.0062821843],"category_scores_gemma":[0.0045712814,0.00042514227,0.00096240273,0.0014422506,0.0009791217,0.0053338115,0.0016050327,0.0022427442,0.00095597096],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016456464,0.00011355333,0.003276086,0.0006646046,0.000092957205,0.00050628075,0.0032729288,0.13113774,0.0041123983,0.3323246,0.024023121,0.50031114],"study_design_scores_gemma":[0.000014964903,0.000022064993,0.0006597994,0.00009633375,0.0000219102,0.00010196446,0.00038172424,0.51401484,0.0014580517,0.47194386,0.0112687545,0.000015752741],"about_ca_topic_score_codex":0.002397991,"about_ca_topic_score_gemma":0.003941934,"teacher_disagreement_score":0.0062821843,"about_ca_system_score_codex":0.00083256565,"about_ca_system_score_gemma":0.0006260984,"threshold_uncertainty_score":0.021016002},"labels":[],"label_agreement":null},{"id":"W4387183694","doi":"10.1111/coin.12603","title":"A semantically enhanced text retrieval framework with abstractive summarization","year":2023,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Barrie Urology Group; York University","funders":"China Scholarship Council; Natural Science Foundation of Hubei Province; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China; Hubei Provincial Department of Education","keywords":"Computer science; Automatic summarization; Natural language processing; Artificial intelligence; Question answering; Encoder; Language model; Information retrieval; Generative grammar; Transformer; Semantics (computer science); Sequence (biology); Programming language","score_opus":0.029918996281097184,"score_gpt":0.299368928543697,"score_spread":0.26944993226259983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387183694","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014681571,0.00085494894,0.97644955,0.00034005655,0.00008233498,0.00012118293,0.00047641038,0.0047336696,0.0022601606],"genre_scores_gemma":[0.49041831,0.0009349749,0.491527,0.00037229128,0.00035982335,0.00028449917,0.004092832,0.0007015927,0.0113085965],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925536,0.00028707457,0.00005619208,0.00014096191,0.00018456028,0.00007578674],"domain_scores_gemma":[0.9993623,0.00022904265,0.000079951096,0.00011089773,0.00018131989,0.00003653408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010460702,0.00088961073,0.00088764384,0.0020897593,0.0003928128,0.0012528265,0.0013229077,0.0008495528,0.003778058],"category_scores_gemma":[0.0024971766,0.00029751082,0.00092734786,0.0015868797,0.0004681732,0.00229054,0.001090232,0.0011155679,0.0022618198],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044172522,0.00034420402,0.00076643925,0.00058432034,0.0001554721,0.0004355703,0.00050141243,0.18376642,0.04491758,0.06160751,0.020978007,0.68550134],"study_design_scores_gemma":[0.00004129252,0.00014620926,0.00019601716,0.000020029456,0.00006986116,0.000105241714,0.00007509245,0.9583941,0.009748288,0.024728958,0.006444154,0.0000306466],"about_ca_topic_score_codex":0.0028809619,"about_ca_topic_score_gemma":0.003750714,"teacher_disagreement_score":0.003778058,"about_ca_system_score_codex":0.0006817575,"about_ca_system_score_gemma":0.0009292898,"threshold_uncertainty_score":0.0126389265},"labels":[],"label_agreement":null},{"id":"W4387185473","doi":"10.3233/faia230559","title":"Investigating the Learning Behaviour of In-Context Learning: A Comparison with Supervised Learning","year":2023,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Dalhousie University; Vector Institute; Western University","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Context (archaeology); Computer science; Task (project management); Artificial intelligence; Machine learning; Cognitive psychology; Psychology; Geography; Engineering","score_opus":0.07672128964706759,"score_gpt":0.2897408337165567,"score_spread":0.2130195440694891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387185473","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20202678,0.014376088,0.74674165,0.0027970993,0.00038457566,0.00028504152,0.00082123664,0.008805263,0.023762327],"genre_scores_gemma":[0.6424283,0.0028623275,0.3437886,0.00071217545,0.0002458415,0.00025556647,0.0019699654,0.0009697901,0.006767554],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970517,0.001635017,0.0001044632,0.0006401173,0.00047317328,0.000095541014],"domain_scores_gemma":[0.98385805,0.011416087,0.0005042062,0.002687818,0.0012614126,0.00027243563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005291657,0.0017056186,0.001057484,0.0008433179,0.00043457333,0.0014867194,0.00212142,0.0015063193,0.0025464382],"category_scores_gemma":[0.018348498,0.0004885531,0.00072205707,0.0011676402,0.000988263,0.004494239,0.0019447893,0.003303449,0.0015266133],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058626925,0.0004964859,0.011373842,0.0009708576,0.00036770134,0.00010162747,0.00057281606,0.23645055,0.0108264815,0.013579271,0.011689053,0.712985],"study_design_scores_gemma":[0.00002491242,0.0003910679,0.0014388973,0.000052817562,0.00004881678,0.000080579266,0.00011855494,0.97188085,0.009441325,0.013159293,0.0033365353,0.0000263029],"about_ca_topic_score_codex":0.003679963,"about_ca_topic_score_gemma":0.0058070533,"teacher_disagreement_score":0.005291657,"about_ca_system_score_codex":0.0011283449,"about_ca_system_score_gemma":0.0009121447,"threshold_uncertainty_score":0.027985275},"labels":[],"label_agreement":null},{"id":"W4387389810","doi":"10.48550/arxiv.2310.02457","title":"The Empty Signifier Problem: Towards Clearer Paradigms for Operationalising \"Alignment\" in Large Language Models","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Economic and Social Research Council; European Commission; University of Oxford; York University","keywords":"Parallels; Transparency (behavior); Through-the-lens metering; Politics; Vocabulary; Epistemology; Sociology; Computer science; Cognitive science; Linguistics; Lens (geology); Political science; Psychology; Economics; Engineering; Law; Philosophy","score_opus":0.12247549650766784,"score_gpt":0.2319155566154391,"score_spread":0.10944006010777127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387389810","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009662534,0.0007499312,0.9762325,0.007568828,0.00016315436,0.000117799675,0.00016945561,0.00024092816,0.0050948467],"genre_scores_gemma":[0.42320335,0.0009724947,0.5697346,0.0019074271,0.00058452744,0.0011426699,0.00071737455,0.0006090201,0.0011285413],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.901719,0.075811796,0.005056969,0.009178881,0.007247352,0.0009859562],"domain_scores_gemma":[0.726062,0.2103182,0.015836013,0.03749309,0.007736514,0.0025542732],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08592337,0.0019556344,0.0027893493,0.007316244,0.00491072,0.018896153,0.0050766785,0.0052563623,0.0051648538],"category_scores_gemma":[0.29611325,0.0015062512,0.0022241063,0.009174901,0.036646977,0.061642013,0.021248437,0.017424693,0.001081315],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000056582783,0.000019155996,0.0018163563,0.00020928813,0.0000543882,0.000054741635,0.007484775,0.0020664106,0.00033872086,0.9692068,0.0007013294,0.017991407],"study_design_scores_gemma":[0.000014166233,0.000023256705,0.00029510455,0.00015323247,0.000020855607,0.00006860159,0.0012713592,0.011942084,0.00033111864,0.97972506,0.0061203004,0.000034841407],"about_ca_topic_score_codex":0.003025432,"about_ca_topic_score_gemma":0.0022929048,"teacher_disagreement_score":0.08592337,"about_ca_system_score_codex":0.0042518442,"about_ca_system_score_gemma":0.0059014913,"threshold_uncertainty_score":0.4544117},"labels":[],"label_agreement":null},{"id":"W4387424400","doi":"10.1016/j.asoc.2023.110901","title":"BERT models for Brazilian Portuguese: Pretraining, evaluation and tokenization analysis","year":2023,"lang":"en","type":"article","venue":"Applied Soft Computing","topic":"Topic Modeling","field":"Computer Science","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Conselho Nacional de Desenvolvimento Científico e Tecnológico","keywords":"Computer science; Artificial intelligence; Natural language processing; Language model; Transformer; Lexical analysis; Sentence; Encoder; Transfer of learning; Portuguese; Textual entailment; Machine translation; Logical consequence; Linguistics","score_opus":0.04187459156628819,"score_gpt":0.3029099270173801,"score_spread":0.2610353354510919,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387424400","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6336412,0.00533876,0.2803812,0.0042946674,0.0020062,0.0007483124,0.019853957,0.026096947,0.02763876],"genre_scores_gemma":[0.87739676,0.0014340068,0.08012144,0.00027712563,0.00020680785,0.0005529448,0.026353383,0.0022277134,0.01142977],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9987747,0.0006287581,0.0000800213,0.00029233415,0.00010601171,0.000118243406],"domain_scores_gemma":[0.9912096,0.0072478713,0.00013329979,0.00038646796,0.0007844425,0.0002382341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035228913,0.0022378871,0.0011615492,0.0014828638,0.0013571616,0.0020695785,0.0019196029,0.0013354251,0.009205168],"category_scores_gemma":[0.013860546,0.0009409004,0.0013301763,0.0011882484,0.0004130099,0.0033549564,0.0012708779,0.0033265688,0.0057028946],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004532875,0.0012748073,0.01923747,0.0016768936,0.00061617326,0.0007589031,0.0032279752,0.35250986,0.0071030124,0.009879888,0.06708401,0.5320982],"study_design_scores_gemma":[0.00013852779,0.00024763003,0.0059318924,0.00023416828,0.00027465777,0.00015419692,0.0009208643,0.96392083,0.0049014855,0.0073286192,0.015857227,0.00008991114],"about_ca_topic_score_codex":0.08541442,"about_ca_topic_score_gemma":0.09154653,"teacher_disagreement_score":0.08541442,"about_ca_system_score_codex":0.0023692122,"about_ca_system_score_gemma":0.0029937283,"threshold_uncertainty_score":0.1698345},"labels":[],"label_agreement":null},{"id":"W4387504521","doi":"10.1016/j.patcog.2023.110037","title":"A topic modeling and image classification framework: The Generalized Dirichlet variational autoencoder","year":2023,"lang":"en","type":"article","venue":"Pattern Recognition","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Latent Dirichlet allocation; Autoencoder; Topic model; Dirichlet distribution; Computer science; Covariance; Inference; Artificial intelligence; Generative model; Latent variable; Representation (politics); Benchmark (surveying); Code (set theory); Pattern recognition (psychology); Machine learning; Artificial neural network; Mathematics; Generative grammar; Statistics; Set (abstract data type); Boundary value problem","score_opus":0.08430624976072619,"score_gpt":0.2919918592378226,"score_spread":0.2076856094770964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387504521","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005340457,0.0010013352,0.9923644,0.0004567257,0.0000740185,0.00002575852,0.00012863107,0.0001189365,0.000489855],"genre_scores_gemma":[0.47898638,0.0049838237,0.50187665,0.00075367506,0.0013333429,0.00044436054,0.0016673846,0.000436166,0.009518226],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982734,0.00086677773,0.00007098177,0.00040218348,0.00025962494,0.00012710739],"domain_scores_gemma":[0.9972076,0.002042975,0.00015030483,0.00022969402,0.0002780854,0.00009136658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034387752,0.0008249883,0.0021248874,0.0016327292,0.0005967052,0.0020433287,0.0025831808,0.0021572874,0.0016648315],"category_scores_gemma":[0.007958892,0.0008097038,0.0018334036,0.0021289394,0.0012255014,0.0025144073,0.0016591105,0.002869256,0.000618645],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021477562,0.00017755688,0.0024257384,0.00036325798,0.00043694032,0.00013154262,0.00042879797,0.5153643,0.0050402726,0.21519206,0.00970263,0.25052214],"study_design_scores_gemma":[0.0000073193173,0.000012331605,0.00020647427,0.0000128233905,0.00001990462,0.000025819241,0.00001142062,0.95924836,0.00028415676,0.03905488,0.0011051561,0.000011286371],"about_ca_topic_score_codex":0.00829565,"about_ca_topic_score_gemma":0.007864736,"teacher_disagreement_score":0.00829565,"about_ca_system_score_codex":0.0013252622,"about_ca_system_score_gemma":0.0013006136,"threshold_uncertainty_score":0.018186152},"labels":[],"label_agreement":null},{"id":"W4387522170","doi":"10.1016/j.eswa.2023.122031","title":"A fast local citation recommendation algorithm scalable to multi-topics","year":2023,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Citation; Similarity (geometry); Margin (machine learning); Information retrieval; Scalability; Space (punctuation); Artificial intelligence; Machine learning; Data mining; Data science; World Wide Web","score_opus":0.03651649654463209,"score_gpt":0.29333741079450815,"score_spread":0.2568209142498761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387522170","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04019489,0.0027724914,0.93010515,0.0009006345,0.00063662813,0.00042610467,0.0024604606,0.018795868,0.0037076846],"genre_scores_gemma":[0.13194151,0.0007835917,0.84297824,0.00033123273,0.0007197152,0.00051047764,0.006653156,0.0006752477,0.015406839],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985996,0.00022600094,0.00012484886,0.00038012041,0.0005326094,0.0001368269],"domain_scores_gemma":[0.9965222,0.0011067472,0.00016876806,0.0007569974,0.0011815164,0.0002637801],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012978276,0.0010577257,0.0025578532,0.0043985364,0.0014651108,0.00205652,0.0031112186,0.0021367162,0.0077194255],"category_scores_gemma":[0.0065035946,0.0007592656,0.0015115566,0.0063109454,0.00035897354,0.0025968864,0.001871017,0.0015275307,0.0070578116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054569345,0.00044481535,0.003052322,0.00027032572,0.00036707942,0.00013651916,0.00008189396,0.051632367,0.013941881,0.0034796435,0.047388047,0.8786594],"study_design_scores_gemma":[0.0002030924,0.00008901528,0.00092810864,0.000015294185,0.000112720474,0.00014543526,0.00004480163,0.97916085,0.0050193737,0.007426647,0.0068161893,0.000038431834],"about_ca_topic_score_codex":0.016240288,"about_ca_topic_score_gemma":0.035310417,"teacher_disagreement_score":0.016240288,"about_ca_system_score_codex":0.000924445,"about_ca_system_score_gemma":0.0032212455,"threshold_uncertainty_score":0.03229153},"labels":[],"label_agreement":null},{"id":"W4387528885","doi":"10.1119/perc.2023.pr.bralin","title":"Analysis of student essays in an introductory physics course using natural language processing","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"National Science Foundation","keywords":"Course (navigation); Computer science; Mathematics education; Physics; Psychology; Astronomy","score_opus":0.03381296974521332,"score_gpt":0.3505497077829594,"score_spread":0.3167367380377461,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387528885","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99385613,0.00019289363,0.002887705,0.0001168204,0.00003645906,0.00005149334,0.0019471017,0.00012669856,0.0007847044],"genre_scores_gemma":[0.97838336,0.00014019557,0.010470325,0.000040467454,0.00007706968,0.00014317175,0.008662958,0.000035558114,0.002046884],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99847406,0.00055704685,0.0002165064,0.00029916235,0.00037066897,0.00008256919],"domain_scores_gemma":[0.98533404,0.008856341,0.002130503,0.0005300213,0.0025422615,0.00060691196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016293274,0.0004479177,0.00037118545,0.0036764788,0.00049510796,0.0007822914,0.00031293876,0.00048311628,0.0008429262],"category_scores_gemma":[0.012525902,0.000105122184,0.0003158049,0.002557627,0.00026639787,0.0005268473,0.00061822654,0.00044050443,0.00068945746],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078872254,0.0016070844,0.54643667,0.00085574115,0.00019140294,0.0017206914,0.008402055,0.0057841665,0.033061244,0.0007783732,0.013011632,0.38736227],"study_design_scores_gemma":[0.000061676794,0.0007960893,0.9069146,0.0000969635,0.000067718036,0.0010429207,0.006113939,0.048833422,0.0132323615,0.0012011256,0.0215403,0.00009895791],"about_ca_topic_score_codex":0.0011928026,"about_ca_topic_score_gemma":0.0034613553,"teacher_disagreement_score":0.0036764788,"about_ca_system_score_codex":0.00036811733,"about_ca_system_score_gemma":0.0002998018,"threshold_uncertainty_score":0.008616805},"labels":[],"label_agreement":null},{"id":"W4387586812","doi":"10.2139/ssrn.4600240","title":"Gcsum: Boosting Informativeness and Factual Consistency of Abstractive Summarization with Semantic Graph and Contrastive Learning","year":2023,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Automatic summarization; Boosting (machine learning); Computer science; Natural language processing; Artificial intelligence; Consistency (knowledge bases); Graph; Linguistics; Psychology; Philosophy; Theoretical computer science","score_opus":0.014318081812067191,"score_gpt":0.23393345896519468,"score_spread":0.21961537715312748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387586812","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.076465696,0.0022129007,0.90656185,0.00065685326,0.00024407098,0.00020557325,0.0011784454,0.009503826,0.0029708236],"genre_scores_gemma":[0.60043406,0.0004540189,0.38720977,0.00037206226,0.00047296795,0.00022336375,0.0052720057,0.0013340794,0.0042276215],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981402,0.0006769065,0.00009219549,0.0005975775,0.00037462945,0.000118538956],"domain_scores_gemma":[0.99348855,0.0042460086,0.000383198,0.00093701715,0.0007610971,0.00018416399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031900348,0.0014840206,0.0018029596,0.0038532214,0.0007203546,0.0018096399,0.0021094338,0.0016832615,0.003649326],"category_scores_gemma":[0.012596775,0.0005514648,0.0010344278,0.0022635444,0.0010242391,0.0032512594,0.0022397724,0.0021428159,0.0012509974],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014759209,0.00043102118,0.0035760663,0.00070738076,0.0004633894,0.00012800394,0.0004329969,0.103963345,0.01960118,0.021672674,0.01895342,0.8285945],"study_design_scores_gemma":[0.000111487236,0.00024134327,0.001292027,0.000038579958,0.00013372499,0.000048503854,0.00006155952,0.9490928,0.006722884,0.038952086,0.0032796683,0.000025366051],"about_ca_topic_score_codex":0.0027199814,"about_ca_topic_score_gemma":0.004850268,"teacher_disagreement_score":0.0038532214,"about_ca_system_score_codex":0.0010166072,"about_ca_system_score_gemma":0.0010512833,"threshold_uncertainty_score":0.016870737},"labels":[],"label_agreement":null},{"id":"W4387796700","doi":"10.48550/arxiv.2310.10808","title":"If the Sources Could Talk: Evaluating Large Language Models for Research Assistance in History","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Metadata; Computer science; World Wide Web; Style (visual arts); Information retrieval; Data science; Natural language processing; History","score_opus":0.38480627533143086,"score_gpt":0.31328411518250615,"score_spread":0.07152216014892471,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387796700","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7962341,0.004058543,0.17959535,0.0032601657,0.0002707656,0.00080909947,0.0034932836,0.0060230847,0.0062555494],"genre_scores_gemma":[0.9134143,0.00039975697,0.08013587,0.00029806257,0.00012331532,0.0005127101,0.003781006,0.00021898338,0.001115939],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.991574,0.00687952,0.00035435945,0.00079575373,0.00027832988,0.00011799534],"domain_scores_gemma":[0.9371032,0.05766261,0.0008714825,0.002604056,0.0010647633,0.0006939078],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.019478966,0.0016729358,0.00083870237,0.0024196762,0.0008654275,0.002422476,0.0017253533,0.0024667166,0.0025353236],"category_scores_gemma":[0.049071133,0.0006179202,0.0011760085,0.0016595308,0.0008076413,0.004058898,0.002610011,0.0020731664,0.0011099576],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006495292,0.0018911599,0.03484738,0.0014697205,0.0016162687,0.00030964773,0.004389243,0.4374151,0.004165549,0.0070139915,0.012162084,0.48822457],"study_design_scores_gemma":[0.00015070692,0.00034025224,0.0022244363,0.00006210018,0.00015659494,0.000047905185,0.0005784155,0.9849739,0.0019852559,0.007841278,0.0016010071,0.00003806346],"about_ca_topic_score_codex":0.008008902,"about_ca_topic_score_gemma":0.00888123,"teacher_disagreement_score":0.980521,"about_ca_system_score_codex":0.001643444,"about_ca_system_score_gemma":0.0014104265,"threshold_uncertainty_score":0.10301584},"labels":[],"label_agreement":null},{"id":"W4387848904","doi":"10.1145/3583780.3615026","title":"Mulco: Recognizing Chinese Nested Named Entities through Multiple Scopes","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Scope (computer science); Named-entity recognition; Entity linking; Sequence (biology); Natural language processing; Artificial intelligence; Sequence labeling; Nested set model; Information retrieval; Programming language; Engineering; Relational database; Task (project management)","score_opus":0.04650581625218009,"score_gpt":0.28302877165170154,"score_spread":0.23652295539952145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387848904","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32782695,0.004014071,0.56821185,0.0009013982,0.00059829024,0.0010215465,0.02210426,0.056028076,0.019293472],"genre_scores_gemma":[0.511473,0.0009160852,0.4140592,0.00040303846,0.00016178963,0.00050718413,0.058864616,0.0010518716,0.012563295],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934167,0.00008567423,0.00006160217,0.00034045122,0.0001049979,0.00006568708],"domain_scores_gemma":[0.99847955,0.00043790974,0.00014862807,0.00041915133,0.00037959803,0.00013512465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00089708244,0.001116688,0.00069936214,0.0027582073,0.0008275814,0.0010517722,0.001205813,0.0008549963,0.0033329604],"category_scores_gemma":[0.0024268737,0.00033595524,0.0008753493,0.0014373556,0.00050000334,0.004242596,0.0017870341,0.0008073355,0.0019225843],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006747726,0.0003123439,0.028921226,0.00094722636,0.0002338012,0.0015699612,0.0016816261,0.015799144,0.07649021,0.0114198215,0.07868524,0.7832646],"study_design_scores_gemma":[0.00010702682,0.0003607783,0.028137054,0.00015205833,0.00021755337,0.0016291159,0.0012688766,0.7810725,0.086178616,0.010872251,0.08982121,0.00018282865],"about_ca_topic_score_codex":0.013416743,"about_ca_topic_score_gemma":0.0258084,"teacher_disagreement_score":0.013416743,"about_ca_system_score_codex":0.00070903933,"about_ca_system_score_gemma":0.0014716983,"threshold_uncertainty_score":0.02667731},"labels":[],"label_agreement":null},{"id":"W4387854289","doi":"10.1145/3583780.3615205","title":"Latent Aspect Detection via Backtranslation Augmentation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Computer science; Context (archaeology); Laptop; Vocabulary; SemEval; Natural language processing; Benchmark (surveying); Focus (optics); Artificial intelligence; Data science; Semantics (computer science); Latent semantic analysis; Natural language; World Wide Web; Task (project management); Linguistics","score_opus":0.03564317263585555,"score_gpt":0.2604512646798721,"score_spread":0.22480809204401658,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387854289","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16006283,0.006505678,0.78326994,0.0016421673,0.0012066575,0.0008258186,0.010524966,0.029016517,0.0069453716],"genre_scores_gemma":[0.49207735,0.001650867,0.46781394,0.0011369693,0.0010188301,0.0008195752,0.027750699,0.0016693682,0.0060622836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967361,0.0011357919,0.0003634719,0.0008720153,0.0007413314,0.00015130754],"domain_scores_gemma":[0.99052143,0.0038159802,0.0010457905,0.002195231,0.0022638114,0.00015770624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018832943,0.0018052085,0.0011300826,0.0033371593,0.00059981196,0.0017220057,0.0010559874,0.0008446959,0.0020496529],"category_scores_gemma":[0.013432381,0.00048717376,0.0018294313,0.00270591,0.00086878816,0.0031315177,0.0020783097,0.0017768331,0.003569425],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009487417,0.00041916428,0.015805596,0.0015336787,0.0003267573,0.0009002118,0.002000423,0.00979429,0.10433508,0.005211214,0.047594067,0.8111309],"study_design_scores_gemma":[0.00025327783,0.0007150074,0.018392835,0.00022790395,0.00054242596,0.0034788442,0.0012953083,0.6927775,0.12958977,0.030576881,0.12186591,0.00028430333],"about_ca_topic_score_codex":0.0024511407,"about_ca_topic_score_gemma":0.0060076974,"teacher_disagreement_score":0.0033371593,"about_ca_system_score_codex":0.00055893284,"about_ca_system_score_gemma":0.0013119027,"threshold_uncertainty_score":0.009959936},"labels":[],"label_agreement":null},{"id":"W4387963603","doi":"10.48550/arxiv.2310.16127","title":"Octopus: A Multitask Model and Toolkit for Arabic Natural Language Generation","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Computer science; Python (programming language); Natural language processing; Arabic; Transformer; Artificial intelligence; Language model; Economic shortage; Machine learning; Programming language; Linguistics; Engineering","score_opus":0.13947226149193978,"score_gpt":0.21580148305899904,"score_spread":0.07632922156705926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387963603","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015222063,0.00054584537,0.8060336,0.0008639621,0.0005342104,0.00057157996,0.012464592,0.15581426,0.007949895],"genre_scores_gemma":[0.25214797,0.0006403453,0.67379236,0.0009107254,0.00018914133,0.002429122,0.036205553,0.012567568,0.021117264],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9996195,0.00013607909,0.000029578689,0.00009804831,0.00007495418,0.000041802],"domain_scores_gemma":[0.9989613,0.0005308679,0.000047446778,0.00017677742,0.0001984806,0.00008513406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012621586,0.0012052653,0.0005136455,0.00085418724,0.0005096187,0.0010391155,0.0023826938,0.00087811775,0.015419143],"category_scores_gemma":[0.005273729,0.0006517253,0.0012238509,0.0006413663,0.00034870167,0.0018235814,0.0018944001,0.0025101712,0.009666408],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010707991,0.00036788685,0.0039298264,0.0008645444,0.00031727622,0.0004366561,0.0007090873,0.22824232,0.009118098,0.019339647,0.2844609,0.4511429],"study_design_scores_gemma":[0.000066153974,0.000052416464,0.00036602453,0.000030846746,0.00002210431,0.000101135935,0.000042841206,0.9542869,0.0043450757,0.012109917,0.028542008,0.000034648237],"about_ca_topic_score_codex":0.00883431,"about_ca_topic_score_gemma":0.014734323,"teacher_disagreement_score":0.015419143,"about_ca_system_score_codex":0.0008404131,"about_ca_system_score_gemma":0.0017807786,"threshold_uncertainty_score":0.051582217},"labels":[],"label_agreement":null},{"id":"W4387968110","doi":"10.1145/3581783.3612183","title":"ECENet: Explainable and Context-Enhanced Network for Muti-modal Fact verification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Natural Science Foundation of China","keywords":"Computer science; Modal; Inference; Context (archaeology); Reinforcement learning; Feature (linguistics); Sentence; Artificial intelligence; Process (computing); Feature extraction; Machine learning; Natural language processing; Programming language","score_opus":0.03436481007912221,"score_gpt":0.2648313612036172,"score_spread":0.23046655112449502,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387968110","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062676296,0.0007325637,0.92613006,0.0010816285,0.00015692385,0.00015878362,0.0009388309,0.003521548,0.0046033747],"genre_scores_gemma":[0.7577033,0.00042144812,0.23284203,0.00046262963,0.000102125145,0.00019730019,0.0022943015,0.0001579909,0.00581888],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99972695,0.000055812176,0.000014478945,0.0001047379,0.00006287034,0.000035095516],"domain_scores_gemma":[0.999076,0.00049764215,0.00010656518,0.00014556208,0.00012633248,0.00004785104],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007129054,0.00061602733,0.00036010548,0.0007503926,0.000411007,0.00060365966,0.0012711345,0.0010337121,0.003267988],"category_scores_gemma":[0.0037522588,0.00028202278,0.00067534164,0.00037488778,0.00058001623,0.001955236,0.0012893657,0.0014499631,0.0003898119],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000582019,0.00028987948,0.0072913063,0.00046498238,0.00027651264,0.0010887533,0.00054652576,0.265374,0.0340581,0.07816007,0.017709488,0.5941584],"study_design_scores_gemma":[0.0000167019,0.00003732558,0.00076061295,0.00001908183,0.000042884192,0.000084843436,0.00002606195,0.9534195,0.004826594,0.037904397,0.0028470696,0.0000149295975],"about_ca_topic_score_codex":0.004199234,"about_ca_topic_score_gemma":0.00901396,"teacher_disagreement_score":0.004199234,"about_ca_system_score_codex":0.00087821856,"about_ca_system_score_gemma":0.00084686215,"threshold_uncertainty_score":0.0109324455},"labels":[],"label_agreement":null},{"id":"W4388000226","doi":"10.1016/j.compag.2023.108330","title":"Automated extraction of domain knowledge in the dairy industry","year":2023,"lang":"en","type":"article","venue":"Computers and Electronics in Agriculture","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Knowledge graph; Artificial intelligence; Domain knowledge; Transformer; Knowledge extraction; Parsing; Graph; Deep learning; Natural language processing; Natural language understanding; Machine learning; Natural language; Theoretical computer science; Engineering","score_opus":0.013981378145358514,"score_gpt":0.2657512442648914,"score_spread":0.2517698661195329,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388000226","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54697514,0.013270061,0.37449104,0.0031674546,0.000271578,0.0005445992,0.032777075,0.0057482785,0.022754744],"genre_scores_gemma":[0.79473174,0.004797532,0.15126774,0.00020243613,0.00016624079,0.000259066,0.043008946,0.00018283,0.0053834803],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992543,0.00022158386,0.000081833015,0.00020606138,0.00015289002,0.00008342521],"domain_scores_gemma":[0.9974222,0.0018937088,0.00015766486,0.00014684256,0.00031437117,0.00006524892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070415327,0.0006393492,0.00051340665,0.0050107883,0.0007533905,0.0014419224,0.00067331886,0.0010034842,0.0015060626],"category_scores_gemma":[0.0040633106,0.0003388394,0.00094668777,0.0037792623,0.0002321446,0.0014907959,0.00084363774,0.0007582501,0.0010852143],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005888958,0.0006440586,0.03881457,0.0025157712,0.0003174698,0.0018626982,0.0016909197,0.041739594,0.028250236,0.010124008,0.037722584,0.8357292],"study_design_scores_gemma":[0.00014488527,0.0002750256,0.09457491,0.0007186959,0.0008012225,0.0017579192,0.0033162292,0.6539279,0.049011063,0.024203932,0.17115697,0.00011128568],"about_ca_topic_score_codex":0.01129271,"about_ca_topic_score_gemma":0.016591465,"teacher_disagreement_score":0.01129271,"about_ca_system_score_codex":0.000794197,"about_ca_system_score_gemma":0.0020094088,"threshold_uncertainty_score":0.022453904},"labels":[],"label_agreement":null},{"id":"W4388184947","doi":"10.48550/arxiv.2310.20558","title":"Breaking the Token Barrier: Chunking and Convolution for Efficient Long Text Classification with BERT","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Alliance de recherche numérique du Canada","keywords":"Computer science; Chunking (psychology); Inference; Security token; Artificial intelligence; Benchmark (surveying); Machine learning","score_opus":0.10389894511885368,"score_gpt":0.19970924822964303,"score_spread":0.09581030311078935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388184947","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.067287184,0.00095997885,0.9163858,0.0005914382,0.00023817249,0.00007287907,0.000555287,0.01025325,0.0036560292],"genre_scores_gemma":[0.7499426,0.0005439045,0.235826,0.0004960825,0.00014948263,0.00013731248,0.0020312767,0.00074465544,0.010128772],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99973136,0.000058728692,0.000016637638,0.000099188306,0.000045801855,0.000048209953],"domain_scores_gemma":[0.99918777,0.00035796344,0.000066415414,0.00023452319,0.00009894226,0.000054240503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009354651,0.0009209899,0.0006287112,0.000626028,0.00045294818,0.0009522707,0.0016723305,0.00087586994,0.0034252163],"category_scores_gemma":[0.003133574,0.00044721656,0.00060707296,0.0006926304,0.0007328801,0.0042731296,0.0012696383,0.0019065349,0.002106048],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089459174,0.00023711252,0.0041627744,0.00024816144,0.00013466207,0.0003266307,0.0004929982,0.32765895,0.03383897,0.042337213,0.01957116,0.5700968],"study_design_scores_gemma":[0.0000113009155,0.00004217525,0.00020308339,0.000011461114,0.000013602081,0.000044397715,0.000027996615,0.97358567,0.007128211,0.016818937,0.0021024544,0.000010672628],"about_ca_topic_score_codex":0.006970503,"about_ca_topic_score_gemma":0.01222159,"teacher_disagreement_score":0.006970503,"about_ca_system_score_codex":0.0010720525,"about_ca_system_score_gemma":0.0010934961,"threshold_uncertainty_score":0.013859868},"labels":[],"label_agreement":null},{"id":"W4388212323","doi":"10.1109/issre59848.2023.00052","title":"Resilience Assessment of Large Language Models under Transient Hardware Faults","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Transient (computer programming); Reliability (semiconductor); Resilience (materials science); Power (physics)","score_opus":0.03433932826595434,"score_gpt":0.32593017493957077,"score_spread":0.29159084667361646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388212323","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8872518,0.001120567,0.099931635,0.00087308587,0.00017954299,0.00012692843,0.00089166034,0.0071785664,0.002446136],"genre_scores_gemma":[0.9905861,0.00009143921,0.008260002,0.00007455261,0.0000108547865,0.000043132684,0.00045440887,0.00015046088,0.0003291317],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837625,0.0004964454,0.0001724493,0.00031953215,0.0004117518,0.00022345944],"domain_scores_gemma":[0.9851638,0.010354036,0.0010907986,0.0019417681,0.0010705406,0.00037897986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023560456,0.0011364466,0.00067865016,0.00085723266,0.0005512778,0.00087879994,0.0011538645,0.0012094728,0.0013711237],"category_scores_gemma":[0.022856126,0.00047133406,0.0008710842,0.0003984075,0.0010876816,0.0016355071,0.0012281353,0.0016834707,0.00040835835],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006667473,0.000096504016,0.007439842,0.00023948998,0.00014745673,0.000317629,0.00021405105,0.9563713,0.010804233,0.0014666398,0.0011746677,0.021061495],"study_design_scores_gemma":[0.000017461893,0.00015947428,0.0011013797,0.000016353384,0.000031027263,0.00006019394,0.00007161315,0.9882918,0.008599219,0.0013191656,0.00031779098,0.000014502064],"about_ca_topic_score_codex":0.008248251,"about_ca_topic_score_gemma":0.0051223994,"teacher_disagreement_score":0.008248251,"about_ca_system_score_codex":0.0011912822,"about_ca_system_score_gemma":0.0011240029,"threshold_uncertainty_score":0.016400456},"labels":[],"label_agreement":null},{"id":"W4388306831","doi":"10.23977/acss.2023.070903","title":"Research on NLP Based Automatic Summarization Generation Method for Medical Texts","year":2023,"lang":"en","type":"article","venue":"Advances in Computer Signals and Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Natural language processing; Similarity (geometry); Task (project management); Artificial intelligence; Sentence; Information retrieval; Domain (mathematical analysis); Semantic similarity; Documentation; Deep learning; Generative grammar","score_opus":0.10849512291477016,"score_gpt":0.4202254896988056,"score_spread":0.31173036678403543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388306831","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010299159,0.0007274315,0.9840634,0.00033070918,0.00011160493,0.00020891549,0.00031903782,0.002788719,0.0011510154],"genre_scores_gemma":[0.1321352,0.0009807728,0.8601607,0.00020066311,0.00027693887,0.00042897428,0.002539964,0.0003188529,0.0029578994],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981458,0.0007783978,0.00019171697,0.00043179293,0.00040148542,0.000050879167],"domain_scores_gemma":[0.9955584,0.0025728827,0.000405966,0.00036838834,0.0010189984,0.00007534486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018620321,0.0010040714,0.0007125237,0.00240246,0.00054034346,0.0012018346,0.0010416404,0.000817436,0.0029272456],"category_scores_gemma":[0.008825839,0.00027596604,0.00087679754,0.0018270565,0.00039240072,0.0021068477,0.00054907944,0.00094777474,0.0019761897],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018044078,0.00014673118,0.0013565259,0.00083925127,0.00012985436,0.00026506666,0.00072037097,0.03213268,0.045254935,0.0096224025,0.0073579336,0.90199375],"study_design_scores_gemma":[0.00006444016,0.0004054634,0.0021224509,0.000098399236,0.00017733293,0.0005311261,0.00037734152,0.9043602,0.05475402,0.01536391,0.021676695,0.00006870015],"about_ca_topic_score_codex":0.0012381421,"about_ca_topic_score_gemma":0.0012622903,"teacher_disagreement_score":0.0029272456,"about_ca_system_score_codex":0.0004911674,"about_ca_system_score_gemma":0.0010963185,"threshold_uncertainty_score":0.009847522},"labels":[],"label_agreement":null},{"id":"W4388328907","doi":"10.48550/arxiv.2311.00913","title":"Self-Influence Guided Data Reweighting for Language Model Pre-training","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Novelty; Artificial intelligence; Context (archaeology); Task (project management); Relevance (law); Language model; Machine learning; Sample (material); Training set; Point (geometry); Stability (learning theory); Quality (philosophy); Natural language processing; Mathematics","score_opus":0.2694218530808317,"score_gpt":0.26253805896662474,"score_spread":0.006883794114206976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388328907","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035708044,0.0008292208,0.9568722,0.0003855106,0.0001805133,0.00015016786,0.00020980126,0.0044938778,0.001170692],"genre_scores_gemma":[0.49109003,0.00045932073,0.49761638,0.0006567492,0.00041673784,0.0006419016,0.0022175203,0.0022115405,0.004689817],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967554,0.0014193841,0.00022760361,0.0007839564,0.0006037286,0.00020999127],"domain_scores_gemma":[0.98808265,0.0070240963,0.0005719938,0.0023935323,0.0015153299,0.0004124239],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060032057,0.0019373898,0.0018339133,0.0014063724,0.000970129,0.002011497,0.0028277526,0.0022208653,0.0027686264],"category_scores_gemma":[0.031033536,0.0010072303,0.0014501999,0.0010344584,0.0016126527,0.003818194,0.0036945743,0.0053691166,0.002229148],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010935501,0.0006857196,0.009006938,0.00064687966,0.00052957766,0.0002970326,0.0010289975,0.29803234,0.03775102,0.014757776,0.01337954,0.62279063],"study_design_scores_gemma":[0.000059033988,0.00016828111,0.0007368437,0.000041025745,0.000050145998,0.00007737008,0.00006688235,0.97317713,0.012576101,0.010390533,0.0026296424,0.00002697525],"about_ca_topic_score_codex":0.0021572649,"about_ca_topic_score_gemma":0.006083444,"teacher_disagreement_score":0.0060032057,"about_ca_system_score_codex":0.0008161211,"about_ca_system_score_gemma":0.0016007401,"threshold_uncertainty_score":0.031748354},"labels":[],"label_agreement":null},{"id":"W4388341202","doi":"10.1007/s10489-023-04992-9","title":"Arabic text detection: a survey of recent progress challenges and opportunities","year":2023,"lang":"en","type":"article","venue":"Applied Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Arabic; Context (archaeology); Representation (politics); Natural language processing; Field (mathematics); Artificial intelligence; Semantics (computer science); Open research; Intermediate language; Data science; Linguistics; World Wide Web; Politics; Programming language; Mathematics","score_opus":0.2370631666371304,"score_gpt":0.31121353525992246,"score_spread":0.07415036862279206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388341202","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021149825,0.7956457,0.13713165,0.016844591,0.0030005267,0.00031049317,0.0016882517,0.004328481,0.019900419],"genre_scores_gemma":[0.17615527,0.6090459,0.17441955,0.00511433,0.009993691,0.0003478935,0.0067066834,0.00076364115,0.017453093],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970504,0.000650479,0.00031248774,0.00074216083,0.0010869605,0.00015745091],"domain_scores_gemma":[0.986858,0.007128385,0.00065323093,0.0005911062,0.004301443,0.0004676824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054422244,0.0013721301,0.0020350658,0.00954142,0.0010316259,0.005042571,0.0021707646,0.0013218593,0.0048170136],"category_scores_gemma":[0.011697716,0.00063181977,0.0008673994,0.0066315914,0.0011387184,0.007585037,0.0016568579,0.0015762412,0.003986047],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010006998,0.000104112,0.0033228218,0.002669928,0.00005350918,0.000057373036,0.00032832625,0.00080505287,0.0030965991,0.0035499972,0.025946243,0.9599659],"study_design_scores_gemma":[0.000060700055,0.00056329643,0.015308288,0.0032594227,0.0005838267,0.002375811,0.004651608,0.08186313,0.026818506,0.030775776,0.83339715,0.00034256195],"about_ca_topic_score_codex":0.003145903,"about_ca_topic_score_gemma":0.0028708528,"teacher_disagreement_score":0.00954142,"about_ca_system_score_codex":0.0009882485,"about_ca_system_score_gemma":0.0019867406,"threshold_uncertainty_score":0.028781533},"labels":[],"label_agreement":null},{"id":"W4388380341","doi":"10.31234/osf.io/dc6tz","title":"Evaluating Large Language Models for Assisting in Meta-Analysis","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Natural Science Foundation of China","keywords":"Coding (social sciences); Meta-analysis; Computer science; Qualitative analysis; Perspective (graphical); Empirical research; Qualitative property; Quantitative analysis (chemistry); Natural language processing; Data science; Qualitative research; Psychology; Artificial intelligence; Machine learning; Statistics; Social science; Sociology; Medicine; Pathology","score_opus":0.48496812663430605,"score_gpt":0.4559516276934493,"score_spread":0.029016498940856728,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388380341","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026896136,0.022490136,0.9174391,0.005374704,0.00066675723,0.006145963,0.008445925,0.010623622,0.0019177016],"genre_scores_gemma":[0.10585295,0.0016685443,0.8803373,0.0006129828,0.00015095541,0.0072815134,0.0031860278,0.0006963821,0.00021339733],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.5984443,0.3594617,0.020916728,0.012171903,0.00845223,0.0005532147],"domain_scores_gemma":[0.10892054,0.85447496,0.011510179,0.018929983,0.0051982887,0.00096607505],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.3893289,0.0065270993,0.006219397,0.020330284,0.003076193,0.012994967,0.0054382123,0.0042382814,0.0054376014],"category_scores_gemma":[0.65858364,0.004208228,0.022635618,0.016154684,0.0023627046,0.011457715,0.007717223,0.0058369855,0.0009850628],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008604903,0.0007570895,0.054494936,0.06606851,0.17962068,0.0009570966,0.013340689,0.14174905,0.0047837626,0.041242376,0.019713808,0.46866712],"study_design_scores_gemma":[0.004110311,0.002149534,0.008422072,0.009504878,0.08720462,0.0006467044,0.0025677406,0.65619045,0.007255632,0.19130936,0.02969033,0.0009483921],"about_ca_topic_score_codex":0.0052123265,"about_ca_topic_score_gemma":0.0147706,"teacher_disagreement_score":0.6106711,"about_ca_system_score_codex":0.0057239905,"about_ca_system_score_gemma":0.011067008,"threshold_uncertainty_score":0.7530662},"labels":[],"label_agreement":null},{"id":"W4388405359","doi":"10.1109/dsaa60987.2023.10302466","title":"Learning Representations through Contrastive Strategies for a more Robust Stance Detection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; Lakehead University","funders":"","keywords":"Computer science; Artificial intelligence; Embedding; Sentence; Transformer; Machine learning; Adversarial system; Feature learning; Natural language processing","score_opus":0.06414987442765853,"score_gpt":0.3212894235715813,"score_spread":0.2571395491439228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388405359","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06533547,0.00031848458,0.9298958,0.00042024892,0.00007908059,0.000087425025,0.00019330389,0.0017889703,0.0018811211],"genre_scores_gemma":[0.76354116,0.00024035407,0.2308738,0.0004852083,0.00011426048,0.00014459535,0.0009409995,0.0002603291,0.003399406],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99924624,0.00026763917,0.000041276173,0.00025445293,0.00012231068,0.0000680025],"domain_scores_gemma":[0.99817014,0.00089762895,0.00021384546,0.00041597674,0.00022644212,0.00007595248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015938856,0.0009820777,0.0006206911,0.0009175429,0.00032076612,0.0010173188,0.0010696898,0.001077265,0.0023163059],"category_scores_gemma":[0.0062023825,0.000302464,0.0007427336,0.00054793607,0.0008903604,0.0028343864,0.0015453574,0.0023605477,0.0015317248],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006744748,0.00066268665,0.00600335,0.00025626592,0.00018409271,0.00045433472,0.0006880312,0.20202422,0.10038999,0.041199554,0.009825682,0.6376374],"study_design_scores_gemma":[0.000022940814,0.00015264377,0.00042655104,0.000019822783,0.000025002106,0.00009223221,0.0000631226,0.9612849,0.01227859,0.023845533,0.0017726485,0.000015991187],"about_ca_topic_score_codex":0.0010453576,"about_ca_topic_score_gemma":0.0018751445,"teacher_disagreement_score":0.0023163059,"about_ca_system_score_codex":0.0006331375,"about_ca_system_score_gemma":0.000619047,"threshold_uncertainty_score":0.008429408},"labels":[],"label_agreement":null},{"id":"W4388477306","doi":"10.18280/ria.370524","title":"Assessing Semantic Similarity Measures and Proposing a WuP-Resnik Hybrid Metric for Enhanced Arabic Language Processing","year":2023,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Arabic; Semantic similarity; Similarity (geometry); Metric (unit); Natural language processing; Computer science; Artificial intelligence; Linguistics; Philosophy; Engineering; Image (mathematics)","score_opus":0.07493175031527047,"score_gpt":0.3316384455456369,"score_spread":0.2567066952303664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388477306","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037580986,0.0007260932,0.957394,0.00022168976,0.000112549795,0.00018235129,0.00021891714,0.0005017337,0.003061655],"genre_scores_gemma":[0.33935255,0.00062378985,0.6571163,0.0000839423,0.0001400788,0.00028543608,0.0007093559,0.0001291668,0.0015594977],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9959269,0.0011225523,0.0005022443,0.00052791776,0.0017202402,0.00020014208],"domain_scores_gemma":[0.99294704,0.0028412614,0.00088973186,0.0008599076,0.002236818,0.00022520783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004196886,0.0011784284,0.0011996283,0.008484272,0.0008542191,0.0030239983,0.0012333414,0.0014108911,0.0015382671],"category_scores_gemma":[0.01980176,0.0002684621,0.0009395947,0.005212631,0.0013702981,0.0058136913,0.0018440912,0.0010948965,0.0010823858],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004256808,0.00024834787,0.007989005,0.00083800376,0.00026757582,0.0003636086,0.00092887454,0.05985392,0.056679048,0.05891406,0.0051147044,0.8083772],"study_design_scores_gemma":[0.000029790897,0.0006718532,0.010314407,0.00010081139,0.00014317328,0.0010453027,0.00073038467,0.84727937,0.0563639,0.070238486,0.012887053,0.00019545296],"about_ca_topic_score_codex":0.0015086167,"about_ca_topic_score_gemma":0.0016489727,"teacher_disagreement_score":0.008484272,"about_ca_system_score_codex":0.00093805796,"about_ca_system_score_gemma":0.0012746729,"threshold_uncertainty_score":0.022195518},"labels":[],"label_agreement":null},{"id":"W4388477679","doi":"10.18280/ria.370513","title":"Enhancing Text Summarization with a T5 Model and Bayesian Optimization","year":2023,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Bayesian optimization; Computer science; Bayesian probability; Artificial intelligence; Information retrieval; Natural language processing","score_opus":0.03066239276200701,"score_gpt":0.2504901021751001,"score_spread":0.2198277094130931,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388477679","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017045943,0.00064870453,0.97676533,0.0005070076,0.00007835695,0.0001769196,0.0002371848,0.002153106,0.0023875106],"genre_scores_gemma":[0.2890221,0.0006903909,0.69947875,0.00045811135,0.0002115381,0.0005011732,0.0017996214,0.00059087586,0.0072474685],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983884,0.00066532585,0.00013046572,0.00033035522,0.00037267516,0.000112850525],"domain_scores_gemma":[0.9967535,0.0017904836,0.00023568199,0.00016789514,0.00096900365,0.00008348702],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028497633,0.001212649,0.00131602,0.0017960706,0.000664994,0.002022082,0.001405387,0.0017410507,0.0023715578],"category_scores_gemma":[0.008422908,0.00061522727,0.0015245171,0.0014642453,0.00047209152,0.0023594026,0.0010409139,0.0016353887,0.0014531569],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035115983,0.00029036673,0.0015108234,0.000345685,0.00022952438,0.00011771361,0.00042066426,0.58287144,0.009678149,0.008322912,0.008701068,0.38716048],"study_design_scores_gemma":[0.000013136576,0.000061007526,0.00021694676,0.000012291375,0.00003141135,0.00001772152,0.000035312893,0.99431545,0.0015412407,0.0023829148,0.0013599339,0.000012652263],"about_ca_topic_score_codex":0.01721243,"about_ca_topic_score_gemma":0.017974196,"teacher_disagreement_score":0.01721243,"about_ca_system_score_codex":0.0015770235,"about_ca_system_score_gemma":0.0016851777,"threshold_uncertainty_score":0.03422445},"labels":[],"label_agreement":null},{"id":"W4388481814","doi":"10.48550/arxiv.2311.02891","title":"AdaFlood: Adaptive Flood Regularization","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Regularization (linguistics); Flood myth; Computer science; Sample (material); Generalization; Artificial intelligence; Artificial neural network; Asynchronous communication; Machine learning; Mathematics; Geography","score_opus":0.12663066713815258,"score_gpt":0.18398905481882552,"score_spread":0.057358387680672945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388481814","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013829808,0.00055953296,0.97882515,0.00036839445,0.00012903569,0.00008463756,0.00022139425,0.0046805046,0.0013016028],"genre_scores_gemma":[0.35589144,0.00051665714,0.6277796,0.0011339774,0.0003069507,0.0006633971,0.001628256,0.0014460032,0.010633707],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99910897,0.00025484271,0.000048432485,0.00023804455,0.00025838552,0.00009127117],"domain_scores_gemma":[0.99870193,0.0005443522,0.00011220847,0.0002451457,0.00031921658,0.000077277975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022612787,0.0014305578,0.001221943,0.0009637314,0.0004953185,0.0011950077,0.0025291485,0.0021257317,0.0024634195],"category_scores_gemma":[0.0055862474,0.00061244576,0.0009301224,0.0006501473,0.0009349581,0.0018984998,0.002088331,0.00295042,0.0013852887],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005246432,0.00037777465,0.002631431,0.00028658457,0.00026875135,0.00015876546,0.00023865381,0.35640648,0.012021528,0.0121759735,0.03313415,0.5817752],"study_design_scores_gemma":[0.00002981032,0.00007019667,0.00022188832,0.00001723508,0.000011113603,0.0000524778,0.00001560144,0.9846884,0.0034889758,0.007816154,0.0035731592,0.000015079182],"about_ca_topic_score_codex":0.0028960828,"about_ca_topic_score_gemma":0.0045997174,"teacher_disagreement_score":0.0028960828,"about_ca_system_score_codex":0.00074819947,"about_ca_system_score_gemma":0.0011490067,"threshold_uncertainty_score":0.011958957},"labels":[],"label_agreement":null},{"id":"W4388555939","doi":"10.48550/arxiv.2311.04900","title":"How Abstract Is Linguistic Generalization in Large Language Models? Experiments with Argument Structure","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Northwestern University; University of Toronto; Stony Brook University; National Science Foundation; Yale University; Heinrich-Heine-Universität Düsseldorf","keywords":"Argument (complex analysis); Verb; Word order; Generalization; Linguistics; Object (grammar); Computer science; Subject (documents); Noun; Natural language processing; Space (punctuation); Artificial intelligence; Mathematics; Philosophy","score_opus":0.07583491214773017,"score_gpt":0.21339672845317248,"score_spread":0.1375618163054423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388555939","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25688264,0.0026652557,0.7149662,0.009107914,0.00034252633,0.00022803457,0.00169709,0.0063590733,0.007751297],"genre_scores_gemma":[0.8795125,0.0011351757,0.11249052,0.0014259416,0.0002153161,0.0002957247,0.002460736,0.00069015974,0.0017739369],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99185014,0.0050404486,0.00039003557,0.0018819625,0.00056880363,0.00026853348],"domain_scores_gemma":[0.9496899,0.03672923,0.001280025,0.010364827,0.0012479009,0.0006881544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016620332,0.0021000986,0.001982182,0.0008117553,0.0007095568,0.003161466,0.002729902,0.0021757912,0.002629267],"category_scores_gemma":[0.06831369,0.0011150914,0.0018062501,0.0010819136,0.0022360575,0.015101651,0.0030560177,0.00544549,0.0014859077],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011496536,0.0003943464,0.02474994,0.000829176,0.0014944053,0.00030360933,0.0019320265,0.6259869,0.0065677534,0.048489016,0.014373933,0.2737292],"study_design_scores_gemma":[0.000098494,0.00018614065,0.0020949284,0.000099854595,0.00013535282,0.00008549313,0.00027417796,0.8185283,0.002810474,0.17284493,0.0027911875,0.000050681076],"about_ca_topic_score_codex":0.0060256063,"about_ca_topic_score_gemma":0.0096595595,"teacher_disagreement_score":0.016620332,"about_ca_system_score_codex":0.0016491847,"about_ca_system_score_gemma":0.0014482037,"threshold_uncertainty_score":0.08789778},"labels":[],"label_agreement":null},{"id":"W4388652381","doi":"","title":"Literary Natural Language Generation with Psychological Traits","year":2020,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Natural language generation; Computer science; Natural (archaeology); Natural language; Natural language processing; Linguistics; Psychology; History; Philosophy","score_opus":0.025977568990003804,"score_gpt":0.25486348771873263,"score_spread":0.22888591872872882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388652381","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73605525,0.0004927454,0.21868263,0.0028543943,0.00017646252,0.00014285289,0.0007461096,0.0014513801,0.039398212],"genre_scores_gemma":[0.98941004,0.000058390604,0.007112984,0.00005376825,0.00005788827,0.00003711441,0.00027173437,0.00013060035,0.002867413],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99801034,0.0012277256,0.000061329374,0.00039980392,0.00020052519,0.00010045589],"domain_scores_gemma":[0.97259414,0.021629412,0.0014561756,0.0013843564,0.0019540829,0.0009819118],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002408239,0.00037851933,0.00035424987,0.0011740788,0.00067222887,0.004027802,0.0004818954,0.00080725126,0.00960257],"category_scores_gemma":[0.028812967,0.0004932287,0.00052844035,0.0008960856,0.0010726374,0.0027411073,0.0014475742,0.0012016166,0.0014484877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022078701,0.0011474189,0.14795314,0.0011387491,0.00049921236,0.003560134,0.022930996,0.05743493,0.05474874,0.34608385,0.019365968,0.342929],"study_design_scores_gemma":[0.00015059365,0.0004875234,0.08866044,0.0001471649,0.00021478582,0.0016619648,0.0050816536,0.6093442,0.010573652,0.2720487,0.01146794,0.00016136043],"about_ca_topic_score_codex":0.0006431279,"about_ca_topic_score_gemma":0.00048600585,"teacher_disagreement_score":0.00960257,"about_ca_system_score_codex":0.0007682212,"about_ca_system_score_gemma":0.000387261,"threshold_uncertainty_score":0.032123804},"labels":[],"label_agreement":null},{"id":"W4388666107","doi":"10.26615/978-954-452-092-2_018","title":"Beyond Information: Is ChatGPT Empathetic Enough?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Classifier (UML); Emotion detection; Emotional intelligence; Human–computer interaction; Artificial intelligence; Emotion recognition; Cognitive psychology; Natural language processing; Psychology; Social psychology","score_opus":0.017049252158564815,"score_gpt":0.23769111733953036,"score_spread":0.22064186518096554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388666107","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6791535,0.0010009692,0.27104986,0.0040131584,0.0006155461,0.000586927,0.0021978102,0.013252729,0.02812945],"genre_scores_gemma":[0.9572977,0.00020081215,0.034108125,0.00073290896,0.00013433033,0.00021790348,0.0014994286,0.00046052324,0.0053482517],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971576,0.0017027315,0.00012928639,0.0004601845,0.0003569453,0.00019324207],"domain_scores_gemma":[0.98795205,0.007136951,0.00068053865,0.0021011513,0.0012125766,0.0009167641],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036571731,0.00057314086,0.0005565977,0.00032932218,0.00059135375,0.0023194442,0.0009185532,0.0011962331,0.00417173],"category_scores_gemma":[0.026537951,0.00023804302,0.00037006842,0.0003138282,0.0006622558,0.0043214,0.0022584102,0.001487021,0.0022666857],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0061617675,0.0010010347,0.080907494,0.0046051336,0.00043017569,0.0030162875,0.036037806,0.017013129,0.12850907,0.02361912,0.044991482,0.65370744],"study_design_scores_gemma":[0.00054081413,0.0051982244,0.12065261,0.0010745024,0.0007631158,0.006631912,0.025785754,0.4210968,0.07519452,0.07882771,0.26365834,0.0005756934],"about_ca_topic_score_codex":0.00062054145,"about_ca_topic_score_gemma":0.0005924205,"teacher_disagreement_score":0.00417173,"about_ca_system_score_codex":0.00040652134,"about_ca_system_score_gemma":0.00046947898,"threshold_uncertainty_score":0.01934123},"labels":[],"label_agreement":null},{"id":"W4388779961","doi":"10.1515/9781552384039-013","title":"Schema-Independent Retrieval from Heterogeneous Structured Text","year":2006,"lang":"en","type":"book-chapter","venue":"University of Calgary Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Schema (genetic algorithms); Information retrieval; Computer science; Natural language processing","score_opus":0.018682485011425987,"score_gpt":0.18439956649265848,"score_spread":0.1657170814812325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388779961","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04646196,0.0056379577,0.90584975,0.0010407554,0.00027171016,0.00037851188,0.0035513015,0.014526313,0.022281682],"genre_scores_gemma":[0.21629775,0.0055704056,0.70800894,0.0005705539,0.0003076531,0.00046996103,0.02672706,0.0025952465,0.039452467],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99907196,0.00026684886,0.000098951736,0.0001473359,0.00034301786,0.00007189534],"domain_scores_gemma":[0.9977984,0.0010819403,0.0000591442,0.0005305291,0.00047550676,0.000054451233],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014930214,0.00062924536,0.0012883235,0.0025640114,0.0005284582,0.0027882154,0.0015365916,0.0008209872,0.008471762],"category_scores_gemma":[0.0043574744,0.0008179379,0.0009870285,0.0037028436,0.00056476216,0.004976387,0.0017905317,0.0010017272,0.008046406],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005922921,0.000181053,0.0009237085,0.0008595245,0.00015821721,0.0002618131,0.00085083034,0.012378858,0.07481625,0.024214173,0.08033938,0.80442387],"study_design_scores_gemma":[0.0003493498,0.00027030258,0.0042486037,0.0002514545,0.00055566814,0.0014990076,0.0012991882,0.57044107,0.16656299,0.08168667,0.17267966,0.0001561111],"about_ca_topic_score_codex":0.0025075818,"about_ca_topic_score_gemma":0.0032397474,"teacher_disagreement_score":0.008471762,"about_ca_system_score_codex":0.00069534447,"about_ca_system_score_gemma":0.0010504859,"threshold_uncertainty_score":0.028340876},"labels":[],"label_agreement":null},{"id":"W4388963827","doi":"10.48550/arxiv.2311.13538","title":"AlignedCoT: Prompting Large Language Models via Native-Speaking Demonstrations","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Guangzhou Municipal Science and Technology Bureau; National Natural Science Foundation of China; Guangdong Science and Technology Department; Institute for Catastrophic Loss Reduction","keywords":"Context (archaeology); Style (visual arts); Computer science; Set (abstract data type); Natural language processing; Code (set theory); Artificial intelligence; Linguistics; Programming language; Geography","score_opus":0.11570230983264429,"score_gpt":0.22074737020777532,"score_spread":0.10504506037513103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388963827","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08404733,0.0009198171,0.7437723,0.0006932742,0.000475937,0.0007054638,0.005813858,0.15363368,0.009938369],"genre_scores_gemma":[0.3979266,0.00033982418,0.5688199,0.0006544766,0.0001168315,0.00088729424,0.01748555,0.004496839,0.009272742],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983431,0.0007679732,0.000076757016,0.00051417376,0.00021758016,0.00008041711],"domain_scores_gemma":[0.9952366,0.0029175524,0.00017945118,0.0011436407,0.0003276214,0.0001951643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017446323,0.0017173421,0.0006544973,0.0004246698,0.0004287689,0.0010812089,0.0023559167,0.0013199596,0.015928635],"category_scores_gemma":[0.016408771,0.00049209065,0.00083274493,0.00037042738,0.0006375548,0.0028755604,0.0032594395,0.002484338,0.0069690635],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018461233,0.0008012157,0.0043517374,0.0018170434,0.00014452277,0.0013522761,0.0028223833,0.062137716,0.076717176,0.0145051,0.12772459,0.7057801],"study_design_scores_gemma":[0.00059461547,0.00092326984,0.0021799803,0.0001718254,0.0000818347,0.0007910925,0.00092674245,0.7746803,0.06279076,0.03953299,0.11715768,0.00016881741],"about_ca_topic_score_codex":0.0018786186,"about_ca_topic_score_gemma":0.0037228207,"teacher_disagreement_score":0.015928635,"about_ca_system_score_codex":0.0005593966,"about_ca_system_score_gemma":0.00084242196,"threshold_uncertainty_score":0.053286552},"labels":[],"label_agreement":null},{"id":"W4388975825","doi":"10.1007/978-3-031-45043-3_4","title":"Text to Data","year":2023,"lang":"en","type":"book-chapter","venue":"Synthesis lectures on data management","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Parsing; Artificial intelligence; Principle of compositionality; Semantic computing; Graph; Embedding; Information retrieval; Semantic Web; Theoretical computer science","score_opus":0.13425949397111772,"score_gpt":0.2972451479885176,"score_spread":0.16298565401739987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388975825","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00031079553,0.0011177729,0.003183657,0.0053929216,0.009771827,0.00042142448,0.09057466,0.0055461703,0.8836808],"genre_scores_gemma":[0.0011145632,0.0006828764,0.0012235893,0.0019414611,0.00094723864,0.00024848556,0.02689243,0.002709998,0.9642394],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9982332,0.00017636477,0.00012023317,0.0003509547,0.00097593846,0.00014338408],"domain_scores_gemma":[0.9931726,0.001421758,0.00025575358,0.0013377537,0.003210416,0.00060172175],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0012307069,0.0014506322,0.0015625883,0.005012567,0.0016503667,0.0094522815,0.002353815,0.00162704,0.8537714],"category_scores_gemma":[0.014573166,0.00076113973,0.000988118,0.006280774,0.0007878461,0.005373264,0.0038391193,0.0027676905,0.8540971],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020442267,0.0000093681965,0.000034720975,0.00013762187,0.0000019179724,0.000009954303,0.000016138594,0.000026893333,0.000094625546,0.0016656416,0.9624136,0.035568964],"study_design_scores_gemma":[0.000007327412,0.0000038105184,0.00010217602,0.00007328122,0.0000017463817,0.000008897094,0.000020479432,0.000028303179,0.000098235294,0.001013128,0.9986388,0.0000038577705],"about_ca_topic_score_codex":0.006282926,"about_ca_topic_score_gemma":0.006240529,"teacher_disagreement_score":0.8537714,"about_ca_system_score_codex":0.0023371384,"about_ca_system_score_gemma":0.003179286,"threshold_uncertainty_score":0.2085774},"labels":[],"label_agreement":null},{"id":"W4388976681","doi":"10.1007/978-3-031-45043-3_6","title":"Data-to-Text","year":2023,"lang":"en","type":"book-chapter","venue":"Synthesis lectures on data management","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; View; Graph database; Information retrieval; Graph; Relational database; Ask price; Database design; Query language; Natural language; Natural language processing; Theoretical computer science","score_opus":0.1522734746914532,"score_gpt":0.29960791980145984,"score_spread":0.14733444511000665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388976681","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00057876884,0.0011483177,0.12134283,0.004632498,0.0041187303,0.00052165875,0.06322752,0.06842936,0.7360004],"genre_scores_gemma":[0.0055644806,0.0013189533,0.030660009,0.0028454121,0.0011199128,0.00050477585,0.060176536,0.034253582,0.86355627],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981614,0.00026821482,0.00016404886,0.00042728905,0.00085827115,0.00012077659],"domain_scores_gemma":[0.9927638,0.0022104955,0.00017826035,0.00262728,0.0018148973,0.00040532672],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0019783888,0.0018101336,0.0013973542,0.0027240564,0.001094435,0.00883363,0.003302669,0.0015556368,0.6889398],"category_scores_gemma":[0.014547537,0.0012195719,0.0013125826,0.0041910307,0.00078285066,0.008739987,0.0051704305,0.0025756338,0.72433686],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000043203545,0.00002143476,0.00006101433,0.00023750168,0.000008713879,0.00003686349,0.00005694808,0.00021677472,0.00072877354,0.009178372,0.85124826,0.13816226],"study_design_scores_gemma":[0.000011544414,0.000005962895,0.00007199463,0.000070191534,0.0000043133837,0.000045221142,0.000027389611,0.0004929458,0.00091180374,0.006527395,0.99182117,0.000009990676],"about_ca_topic_score_codex":0.002112794,"about_ca_topic_score_gemma":0.0020594462,"teacher_disagreement_score":0.6889398,"about_ca_system_score_codex":0.001409206,"about_ca_system_score_gemma":0.0020951768,"threshold_uncertainty_score":0.44368958},"labels":[],"label_agreement":null},{"id":"W4388976713","doi":"10.1007/978-3-031-45043-3_2","title":"Building an NLIDB: The Basics","year":2023,"lang":"en","type":"book-chapter","venue":"Synthesis lectures on data management","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Leverage (statistics); Sect; Construct (python library); Computer science; Baseline (sea); Rest (music); Programming language; Artificial intelligence; Political science","score_opus":0.09428031179532508,"score_gpt":0.2902417105555458,"score_spread":0.19596139876022073,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388976713","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014874443,0.0016130733,0.89794725,0.0025314395,0.0005299159,0.00022902848,0.001947026,0.014481141,0.07923369],"genre_scores_gemma":[0.017220212,0.0020887607,0.8920112,0.00086835626,0.0003133646,0.0005446728,0.006515993,0.0070681944,0.07336929],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981888,0.00042582146,0.00015171083,0.0003672204,0.00074193336,0.00012454069],"domain_scores_gemma":[0.99819297,0.0006523421,0.00004967262,0.0005298953,0.00042035026,0.00015479875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027355277,0.0007780897,0.00095605204,0.0017705534,0.0018627917,0.007484994,0.00229488,0.0013202142,0.058799904],"category_scores_gemma":[0.007882097,0.0011726901,0.0011256814,0.0027395475,0.0014960185,0.012433934,0.0056876736,0.0026841098,0.06308035],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007674133,0.00008820785,0.0006063221,0.00055956363,0.000019795305,0.000102554724,0.0012500678,0.0026741114,0.0047287066,0.20820238,0.14576964,0.63592184],"study_design_scores_gemma":[0.000014744672,0.000015878168,0.0002125669,0.00021868628,0.000015387637,0.00018727362,0.00031540985,0.016203368,0.0045144544,0.12445811,0.85381424,0.000029919145],"about_ca_topic_score_codex":0.0046553886,"about_ca_topic_score_gemma":0.004239787,"teacher_disagreement_score":0.058799904,"about_ca_system_score_codex":0.0015159624,"about_ca_system_score_gemma":0.002016144,"threshold_uncertainty_score":0.19670528},"labels":[],"label_agreement":null},{"id":"W4388989556","doi":"10.1016/j.ipm.2023.103585","title":"Number-enhanced representation with hierarchical recursive tree decoding for math word problem solving","year":2023,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Natural Science Foundation of China","keywords":"Computer science; Decoding methods; Theoretical computer science; Benchmark (surveying); Tree (set theory); Representation (politics); Artificial intelligence; Algorithm; Mathematics","score_opus":0.0228806507692149,"score_gpt":0.28551661684703245,"score_spread":0.26263596607781753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388989556","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017357992,0.00020442749,0.9749414,0.0001744254,0.0000897954,0.00005058488,0.00040833146,0.0023531402,0.0044200155],"genre_scores_gemma":[0.29259282,0.00028777588,0.6979394,0.00013494444,0.000082566374,0.0001929367,0.0018461673,0.00042370264,0.006499786],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948007,0.00018131133,0.000038230883,0.00012534564,0.0000975974,0.00007752508],"domain_scores_gemma":[0.99919957,0.00040978816,0.000040995797,0.00013536506,0.00016978435,0.000044556615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00047525638,0.0005553052,0.0007609896,0.0008627155,0.00044260285,0.0011520507,0.0011679934,0.0007877519,0.007924383],"category_scores_gemma":[0.0032725607,0.00022894764,0.00073099084,0.001248469,0.00040112986,0.002471753,0.0011399948,0.0011106669,0.0031338288],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040504322,0.000244342,0.0006546993,0.00024503458,0.000051763393,0.00012007193,0.00046679672,0.06970283,0.019335145,0.10038883,0.01245553,0.79592985],"study_design_scores_gemma":[0.000024881661,0.00004814736,0.00019985724,0.000019740162,0.000021772697,0.00003821118,0.000061411454,0.9214734,0.0051483875,0.06859022,0.00435385,0.000020141753],"about_ca_topic_score_codex":0.0040905518,"about_ca_topic_score_gemma":0.0052860104,"teacher_disagreement_score":0.007924383,"about_ca_system_score_codex":0.00047039308,"about_ca_system_score_gemma":0.0014533895,"threshold_uncertainty_score":0.026509702},"labels":[],"label_agreement":null},{"id":"W4389241970","doi":"10.1007/978-3-031-48593-0_12","title":"TON-ViT: A Neuro-Symbolic AI Based on Task Oriented Network with a Vision Transformer","year":2023,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Transformer; Artificial intelligence; Graph; Task (project management); Natural language processing; Theoretical computer science","score_opus":0.012252591426379182,"score_gpt":0.23904105022122601,"score_spread":0.22678845879484683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389241970","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007576914,0.00029188197,0.9745871,0.0001499402,0.00023301263,0.00008474854,0.0002387129,0.006258935,0.010578808],"genre_scores_gemma":[0.43707967,0.00050020864,0.53874165,0.00039081206,0.00010686877,0.0001909918,0.0008121294,0.00063835713,0.021539431],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999866,0.000014964419,0.000007662534,0.00004756009,0.00004226344,0.000021516338],"domain_scores_gemma":[0.9998375,0.000050551615,0.000009943662,0.00003107924,0.0000490367,0.000021960663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00024540717,0.0005435919,0.000532365,0.0003323697,0.0003279449,0.0008970534,0.0016347037,0.00069045497,0.00975271],"category_scores_gemma":[0.00067261525,0.00030551723,0.000505538,0.00046486442,0.00041070662,0.0010490932,0.000911227,0.0011623817,0.0017622329],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039978928,0.00020881879,0.00050454575,0.00033747748,0.00020773824,0.0002907077,0.00013335222,0.18181065,0.056980513,0.043917593,0.0197674,0.69544137],"study_design_scores_gemma":[0.000027065542,0.00009289413,0.00016861541,0.000015861015,0.00004403397,0.000092585025,0.00001731405,0.9642188,0.01169503,0.015800048,0.0078124837,0.000015183059],"about_ca_topic_score_codex":0.0067798938,"about_ca_topic_score_gemma":0.007867423,"teacher_disagreement_score":0.00975271,"about_ca_system_score_codex":0.00050835294,"about_ca_system_score_gemma":0.00086111686,"threshold_uncertainty_score":0.032626092},"labels":[],"label_agreement":null},{"id":"W4389363042","doi":"10.48550/arxiv.2312.00949","title":"Hyperparameter Optimization for Large Language Model Instruction-Tuning","year":2023,"lang":"en","type":"preprint","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs; Hydro-Québec","keywords":"Hyperparameter; Computer science; Pipeline (software); Rank (graph theory); Set (abstract data type); Decomposition; Adaptation (eye); Machine learning; Language model; Artificial intelligence; Fine-tuning; Adaptive optimization; Programming language","score_opus":0.025111019535983572,"score_gpt":0.2599844484148164,"score_spread":0.2348734288788328,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389363042","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031915054,0.0013572574,0.9512773,0.000710111,0.00015721496,0.0001508767,0.00028154164,0.009854615,0.0042959587],"genre_scores_gemma":[0.5460009,0.0005069893,0.44085854,0.0012350285,0.00014454058,0.0008452594,0.0012869254,0.0038200973,0.0053017074],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982241,0.0009782917,0.00007932029,0.0003469614,0.00020793147,0.0001633867],"domain_scores_gemma":[0.99509037,0.0035492608,0.00015495291,0.00073746877,0.0003419796,0.00012592033],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031937456,0.0019197263,0.0013124557,0.0009628356,0.00076863274,0.0017582428,0.0021346484,0.0021752072,0.0065249964],"category_scores_gemma":[0.019731835,0.00097082433,0.0010338752,0.00080613693,0.001225421,0.0025113598,0.002415411,0.0044551357,0.0032384836],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030970582,0.00022827067,0.001659657,0.0002800932,0.0001907371,0.00017876463,0.00031552883,0.77721953,0.0073813517,0.015914602,0.012557109,0.18376474],"study_design_scores_gemma":[0.000031438223,0.000019324743,0.00012168458,0.000021140204,0.000012390857,0.000020388798,0.000029669432,0.9808836,0.0017657189,0.015601055,0.0014818378,0.000011747179],"about_ca_topic_score_codex":0.0042709,"about_ca_topic_score_gemma":0.008835417,"teacher_disagreement_score":0.0065249964,"about_ca_system_score_codex":0.0013858876,"about_ca_system_score_gemma":0.0016176221,"threshold_uncertainty_score":0.021828353},"labels":[],"label_agreement":null},{"id":"W4389371257","doi":"10.1016/j.nlp.2023.100046","title":"Context is not key: Detecting Alzheimer’s disease with both classical and transformer-based neural language models","year":2023,"lang":"en","type":"article","venue":"Natural Language Processing Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; Dalhousie University; Toronto Rehabilitation Institute; University of Toronto; University Health Network","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Transformer; Language model; Word2vec; Artificial intelligence; Machine learning; Artificial neural network; Natural language processing; Speech recognition","score_opus":0.028935805657169744,"score_gpt":0.2818778643010476,"score_spread":0.25294205864387787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389371257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76654,0.0033690513,0.21918723,0.0012900305,0.00019240729,0.00007726256,0.0012546412,0.0024991315,0.0055902717],"genre_scores_gemma":[0.9799973,0.00031045318,0.017223312,0.0001533515,0.000041821382,0.00002260311,0.0009764045,0.000031234827,0.001243569],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997365,0.000096590156,0.000014914129,0.00007930921,0.00003312235,0.000039594503],"domain_scores_gemma":[0.99954826,0.00026864692,0.000034074266,0.000039664366,0.00007871796,0.00003052202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00081963016,0.0006039485,0.00036390912,0.00062839437,0.00016944694,0.00053530757,0.0005197416,0.00038779466,0.00059659546],"category_scores_gemma":[0.0017676583,0.00016875313,0.00050928537,0.00039421747,0.0002601113,0.0009079441,0.0005526888,0.0007531629,0.00040031693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014567015,0.00044937074,0.040426463,0.00021284746,0.00028748403,0.00057597377,0.00034976308,0.3784644,0.015453738,0.007786413,0.008581165,0.54595566],"study_design_scores_gemma":[0.000016017368,0.00007524425,0.0023075484,0.000013929702,0.000045437737,0.000086992906,0.000038053644,0.99158317,0.0017571384,0.0034516018,0.00061494525,0.000009996856],"about_ca_topic_score_codex":0.008014279,"about_ca_topic_score_gemma":0.01437728,"teacher_disagreement_score":0.008014279,"about_ca_system_score_codex":0.00040150914,"about_ca_system_score_gemma":0.0006030834,"threshold_uncertainty_score":0.015935302},"labels":[],"label_agreement":null},{"id":"W4389387146","doi":"10.3233/sji-230063","title":"Classifying respondent comments from the 2021 Canadian Census of Population using machine learning methods1","year":2023,"lang":"en","type":"article","venue":"Statistical Journal of the IAOS","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Respondent; Census; Computer science; Categorization; Encoder; Population; Artificial intelligence; Transformer; Machine learning; Natural language processing; Statistics; Econometrics; Geography; Demography; Mathematics; Sociology; Engineering; Political science","score_opus":0.09543310328496614,"score_gpt":0.3544916148562567,"score_spread":0.25905851157129056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389387146","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.78079134,0.00072305964,0.04488508,0.002673218,0.00052208407,0.0044802744,0.13662225,0.0026201238,0.026682612],"genre_scores_gemma":[0.78271973,0.0009136779,0.096260466,0.00062558684,0.0001703096,0.0033287539,0.08965196,0.00040574695,0.025923776],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962606,0.0007879216,0.00023731792,0.00029244547,0.002102459,0.0003192548],"domain_scores_gemma":[0.9703976,0.0069690403,0.0012590055,0.0010978822,0.019500425,0.0007760178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006861367,0.0007208467,0.0003344322,0.0043978305,0.0013920492,0.0009892804,0.00089237705,0.00040333654,0.003775087],"category_scores_gemma":[0.028393129,0.00019025705,0.0004026518,0.004971301,0.0004269146,0.00052607455,0.00097445183,0.00070887484,0.0018769911],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079471176,0.00041906495,0.3859758,0.0017989174,0.00014896704,0.00039567976,0.020685518,0.0082466155,0.013014741,0.0028895221,0.13584408,0.42978644],"study_design_scores_gemma":[0.000099744386,0.00030920873,0.6709774,0.00053840526,0.000117808595,0.00015465324,0.030525718,0.07278279,0.018907186,0.0018546438,0.20343998,0.00029241893],"about_ca_topic_score_codex":0.77178556,"about_ca_topic_score_gemma":0.8593895,"teacher_disagreement_score":0.22821444,"about_ca_system_score_codex":0.0077643706,"about_ca_system_score_gemma":0.015310911,"threshold_uncertainty_score":0.45911688},"labels":[],"label_agreement":null},{"id":"W4389472577","doi":"10.1016/j.eswa.2023.122676","title":"Top2Label: Explainable zero shot topic labelling using knowledge graphs","year":2023,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cape Breton University; Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Natural language processing; Graph; Sentence; Language model; Knowledge graph; Semantic similarity; Information retrieval; Theoretical computer science","score_opus":0.06576893595542228,"score_gpt":0.3079112960256666,"score_spread":0.2421423600702443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389472577","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0058355336,0.000777616,0.8437394,0.0002790982,0.00026810062,0.00028277835,0.02097715,0.12299141,0.0048488663],"genre_scores_gemma":[0.08054455,0.000606363,0.81147665,0.00036208573,0.00017180627,0.0005573466,0.08514921,0.010715685,0.010416256],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985415,0.0002976097,0.0000628984,0.00059857935,0.00033545826,0.00016401576],"domain_scores_gemma":[0.99765825,0.0010630785,0.00009016296,0.000749165,0.00032607015,0.00011324489],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001430656,0.0025390417,0.0014007405,0.0041318056,0.0015019809,0.0032467702,0.003613044,0.003356816,0.026176536],"category_scores_gemma":[0.007411274,0.0013632169,0.0023792202,0.0029559052,0.0006429465,0.0045011695,0.004063828,0.002849583,0.014485311],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010412239,0.00032789822,0.0015280048,0.0015016799,0.000360536,0.00040260077,0.0005965977,0.020479841,0.014152605,0.018675929,0.23270497,0.7082282],"study_design_scores_gemma":[0.00023789342,0.00013796001,0.001498957,0.00031844713,0.00021785364,0.00041777207,0.00034182705,0.7209419,0.026937407,0.11899109,0.12981753,0.00014133997],"about_ca_topic_score_codex":0.01282687,"about_ca_topic_score_gemma":0.027108569,"teacher_disagreement_score":0.026176536,"about_ca_system_score_codex":0.0015774852,"about_ca_system_score_gemma":0.0017863757,"threshold_uncertainty_score":0.08756924},"labels":[],"label_agreement":null},{"id":"W4389518176","doi":"10.18653/v1/2023.blackboxnlp-1.27","title":"Systematic Generalization by Finetuning? Analyzing Pretrained Language Models Using Constituency Tests","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Samsung; Compute Canada","keywords":"Sentence; Computer science; Natural language processing; Artificial intelligence; Generalization; Copying; Replicate; Language model; Word (group theory); Linguistics","score_opus":0.03343015652090418,"score_gpt":0.27762850851349474,"score_spread":0.24419835199259055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518176","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9256384,0.00026487815,0.06987706,0.00041147028,0.000066773406,0.00013395914,0.0002801518,0.0015231699,0.0018042118],"genre_scores_gemma":[0.98697644,0.0000553822,0.011770081,0.00014151393,0.000012398484,0.0001095771,0.0003944689,0.00017755596,0.0003626519],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99813807,0.00083509047,0.0001637234,0.00053699146,0.00018623323,0.00013986013],"domain_scores_gemma":[0.9790243,0.013038787,0.0011186784,0.0056547886,0.0007560353,0.00040747705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004760685,0.0012309327,0.00086822786,0.00037554099,0.00024126777,0.0011974792,0.0014085217,0.0010219048,0.0016221947],"category_scores_gemma":[0.032411665,0.0007066416,0.00077332626,0.00024304808,0.0009455758,0.0035403918,0.0014389977,0.0032084147,0.0004106197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013712513,0.0009706913,0.051359918,0.0005328093,0.0010608453,0.00043807595,0.0014846754,0.61677873,0.11375831,0.004074625,0.0032442785,0.20492582],"study_design_scores_gemma":[0.0000627179,0.0005666709,0.010335396,0.000030096851,0.00012663893,0.00007568657,0.0001923037,0.94664377,0.029296217,0.011703142,0.0009073805,0.00005994905],"about_ca_topic_score_codex":0.0024104363,"about_ca_topic_score_gemma":0.0027423354,"teacher_disagreement_score":0.004760685,"about_ca_system_score_codex":0.0007450063,"about_ca_system_score_gemma":0.0006100066,"threshold_uncertainty_score":0.02517718},"labels":[],"label_agreement":null},{"id":"W4389518188","doi":"10.18653/v1/2023.banglalp-1.10","title":"BanglaCHQ-Summ: An Abstractive Summarization Dataset for Medical Queries in Bangla Conversational Speech","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Queen's University","funders":"","keywords":"Automatic summarization; Computer science; Popularity; Natural language processing; Bengali; Parsing; Artificial intelligence; Information retrieval; Process (computing); Psychology","score_opus":0.036339229865583106,"score_gpt":0.31397856673610597,"score_spread":0.2776393368705229,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518188","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.072521105,0.0021429285,0.0120378295,0.0015695266,0.0005356944,0.0016512449,0.8874777,0.01092367,0.011140277],"genre_scores_gemma":[0.032586604,0.0002250569,0.013596894,0.00023286733,0.0000755106,0.00083372055,0.9490068,0.00019342976,0.0032490878],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974859,0.0007904925,0.00047776094,0.00054313,0.00052405463,0.00017861494],"domain_scores_gemma":[0.9961351,0.0015438345,0.00024068987,0.0005298958,0.0012395793,0.00031087047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001451308,0.001604371,0.00083984033,0.0027681724,0.0014315424,0.0012565085,0.0015995405,0.002075566,0.009911063],"category_scores_gemma":[0.0066227857,0.0003080986,0.0010037922,0.002289259,0.000626375,0.0013434349,0.0018377703,0.0012546147,0.014216035],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014683382,0.0010318299,0.012884652,0.0063116467,0.0002703674,0.0015788472,0.002663017,0.0040131724,0.031652514,0.0018551458,0.7740357,0.16223487],"study_design_scores_gemma":[0.0007218707,0.00090018485,0.090920396,0.0005956612,0.00026609047,0.002864754,0.006187119,0.038404282,0.029451163,0.0031378844,0.82615155,0.00039901506],"about_ca_topic_score_codex":0.01763843,"about_ca_topic_score_gemma":0.028410815,"teacher_disagreement_score":0.01763843,"about_ca_system_score_codex":0.0015963262,"about_ca_system_score_gemma":0.0020375883,"threshold_uncertainty_score":0.03507155},"labels":[],"label_agreement":null},{"id":"W4389518275","doi":"10.18653/v1/2023.blackboxnlp-1.7","title":"Self-Consistency of Large Language Models under Ambiguity","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Open Philanthropy Project; National Science Foundation","keywords":"Consistency (knowledge bases); Ambiguity; Computer science; Causal consistency; Sequential consistency; Nonparametric statistics; Consistency model; Weak consistency; Probability distribution; Robustness (evolution); Econometrics; Strong consistency; Artificial intelligence; Mathematics; Statistics; Algorithm","score_opus":0.03603277234828802,"score_gpt":0.28146139948881355,"score_spread":0.24542862714052552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518275","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7417682,0.00034055757,0.2509092,0.0011156074,0.000053216914,0.0001893817,0.00040486472,0.0028980072,0.0023208898],"genre_scores_gemma":[0.9605045,0.00006153585,0.037727386,0.00025104074,0.000025852707,0.00012678826,0.00060722645,0.00034403772,0.00035164406],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98292404,0.010557984,0.0012519234,0.0029234742,0.0017835273,0.00055907737],"domain_scores_gemma":[0.79018337,0.16303766,0.010274912,0.029924354,0.0044564256,0.002123207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032557305,0.0011001602,0.0014477362,0.0013656319,0.001115342,0.0034003367,0.0025731397,0.001907054,0.0011983156],"category_scores_gemma":[0.16440277,0.0011979709,0.0014562572,0.0008403519,0.0032648058,0.0077733914,0.0038044022,0.0039767907,0.00050031894],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022828924,0.000690733,0.06450624,0.0005403585,0.0011250959,0.000643756,0.0048826006,0.80264837,0.0139163,0.023018561,0.0026059826,0.08313906],"study_design_scores_gemma":[0.00009929952,0.00032208321,0.0036792932,0.000046435587,0.00008254358,0.0001708683,0.0003330506,0.94172496,0.005167729,0.047665626,0.00064644986,0.00006171378],"about_ca_topic_score_codex":0.003224893,"about_ca_topic_score_gemma":0.0038935074,"teacher_disagreement_score":0.032557305,"about_ca_system_score_codex":0.001691639,"about_ca_system_score_gemma":0.002018791,"threshold_uncertainty_score":0.17218155},"labels":[],"label_agreement":null},{"id":"W4389518298","doi":"10.18653/v1/2023.arabicnlp-1.79","title":"GYM at Qur’an QA 2023 Shared Task: Multi-Task Transfer Learning for Quranic Passage Retrieval and Question Answering with Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Question answering; Computer science; Baseline (sea); Task (project management); Natural language processing; Reading comprehension; Artificial intelligence; Set (abstract data type); Sentence; Transfer of learning; Reading (process); Test (biology); Language model; Encoder; Test set; Information retrieval; Linguistics","score_opus":0.028466764380278568,"score_gpt":0.2720181519120629,"score_spread":0.24355138753178435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518298","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26673368,0.005791183,0.5754073,0.004213836,0.0026601597,0.002434595,0.01800334,0.09847244,0.026283383],"genre_scores_gemma":[0.69605005,0.00047547967,0.20957907,0.0014112147,0.0005777455,0.0019729107,0.055235513,0.0018758735,0.03282219],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979923,0.00095645286,0.00008491216,0.0006047771,0.00020019672,0.0001613361],"domain_scores_gemma":[0.99649626,0.0014184794,0.00009646165,0.00093923835,0.00073410297,0.00031545266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055232714,0.002650367,0.0017320457,0.0011743181,0.0015384562,0.0015427918,0.0033094895,0.00354463,0.015614743],"category_scores_gemma":[0.00951558,0.00067330536,0.0016589699,0.0011149708,0.0008333641,0.0040155384,0.0036075774,0.004554708,0.010519701],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015361517,0.002591522,0.0024356733,0.00085649645,0.00056692184,0.00074464706,0.0006818836,0.063525915,0.02557327,0.0038316436,0.17915271,0.7185032],"study_design_scores_gemma":[0.0004705103,0.0008167769,0.0035384896,0.000054747677,0.00012325104,0.00022578756,0.000316436,0.9458795,0.016296482,0.00986692,0.02229282,0.00011825217],"about_ca_topic_score_codex":0.021806393,"about_ca_topic_score_gemma":0.020118045,"teacher_disagreement_score":0.021806393,"about_ca_system_score_codex":0.0017697478,"about_ca_system_score_gemma":0.002777596,"threshold_uncertainty_score":0.052236557},"labels":[],"label_agreement":null},{"id":"W4389518321","doi":"10.18653/v1/2023.arabicnlp-1.25","title":"Arabic Fine-Grained Entity Recognition","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia Hospital","funders":"","keywords":"Computer science; Natural language processing; Named-entity recognition; Arabic; Artificial intelligence; Entity linking; Modern Standard Arabic; Information retrieval; Linguistics; Knowledge base","score_opus":0.0572780891113574,"score_gpt":0.25948640639402515,"score_spread":0.20220831728266775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518321","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19670963,0.0061623654,0.4759842,0.002756874,0.0034992592,0.0009662961,0.10268182,0.12370283,0.087536834],"genre_scores_gemma":[0.36527634,0.0016076909,0.40214998,0.0012459835,0.00027325348,0.0006195872,0.18497363,0.003188805,0.040664732],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99844044,0.0002954951,0.00019828594,0.00064528635,0.00029647964,0.0001240578],"domain_scores_gemma":[0.9959929,0.0010450868,0.00017809622,0.0010682676,0.0015857513,0.00012981438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014703198,0.002018644,0.0008458759,0.0029138569,0.0012337709,0.0016169824,0.0014087463,0.0011405178,0.021975135],"category_scores_gemma":[0.007600793,0.00044002765,0.0008338363,0.0018820292,0.0006103893,0.0045600533,0.002856185,0.0018398033,0.01718617],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008728896,0.00016996976,0.005386044,0.0021940821,0.00014365125,0.001691304,0.0011659296,0.014325024,0.05440681,0.012590188,0.2256479,0.6814062],"study_design_scores_gemma":[0.00014972886,0.00022045059,0.018047621,0.0007173399,0.00019216347,0.002954461,0.0019030172,0.23248807,0.13786179,0.018803792,0.5862534,0.00040810404],"about_ca_topic_score_codex":0.00851645,"about_ca_topic_score_gemma":0.013132113,"teacher_disagreement_score":0.021975135,"about_ca_system_score_codex":0.00090772397,"about_ca_system_score_gemma":0.0011432862,"threshold_uncertainty_score":0.073514104},"labels":[],"label_agreement":null},{"id":"W4389518345","doi":"10.18653/v1/2023.banglalp-1.41","title":"Semantics Squad at BLP-2023 Task 2: Sentiment Analysis of Bangla Text with Fine Tuned Transformer Based Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Bengali; Computer science; Transformer; Natural language processing; Task (project management); Artificial intelligence; Sentiment analysis; Semantics (computer science); Programming language; Engineering","score_opus":0.02272062310646166,"score_gpt":0.23852026414801658,"score_spread":0.2157996410415549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518345","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.62339556,0.004235753,0.12556697,0.00552791,0.00495522,0.0026006862,0.10834214,0.08284899,0.04252683],"genre_scores_gemma":[0.66991377,0.0005039834,0.11159155,0.0013265726,0.000490934,0.0011117404,0.19145733,0.0030487916,0.02055538],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977075,0.0007484289,0.00020874436,0.0007593589,0.00032784333,0.0002482314],"domain_scores_gemma":[0.9960967,0.0011881868,0.00015182488,0.0009350137,0.0012191521,0.00040911406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036381606,0.003631122,0.0012882132,0.0014400402,0.0012182536,0.0022530335,0.0019739799,0.0021150683,0.008873632],"category_scores_gemma":[0.007984844,0.00061658944,0.0025472024,0.0011266033,0.0006229918,0.0037688657,0.0026700231,0.00326151,0.010643519],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003448977,0.0023501692,0.023997186,0.0025022554,0.00095958076,0.0012275518,0.0017188935,0.036006343,0.060992982,0.0034505182,0.46241125,0.4009343],"study_design_scores_gemma":[0.0010226081,0.0018773698,0.030493984,0.00027551068,0.0004719511,0.0012595148,0.0027331605,0.7819407,0.058579978,0.012046224,0.10899275,0.00030629148],"about_ca_topic_score_codex":0.009944326,"about_ca_topic_score_gemma":0.01739917,"teacher_disagreement_score":0.009944326,"about_ca_system_score_codex":0.001583103,"about_ca_system_score_gemma":0.0018704524,"threshold_uncertainty_score":0.029685259},"labels":[],"label_agreement":null},{"id":"W4389518355","doi":"10.18653/v1/2023.banglalp-1.48","title":"BLP-2023 Task 2: Sentiment Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Scripting language; Task (project management); Domain (mathematical analysis); Sentiment analysis; Artificial intelligence; Natural language processing; Data science; Machine learning; Programming language; Systems engineering","score_opus":0.02410934100984339,"score_gpt":0.26563425604701146,"score_spread":0.24152491503716808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518355","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06099931,0.005666263,0.09620082,0.004435454,0.0068150735,0.005487014,0.5913903,0.16166615,0.067339696],"genre_scores_gemma":[0.05744357,0.0005743543,0.07572265,0.001441538,0.0008726318,0.0037952715,0.8299571,0.0058556884,0.024337236],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99533784,0.0012602392,0.00033817042,0.0014192661,0.0010579235,0.00058655435],"domain_scores_gemma":[0.9955059,0.0010464618,0.0002505218,0.0011046985,0.0014899762,0.00060247805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004635052,0.0057091434,0.0024573046,0.003225235,0.0024952414,0.0040185354,0.0027307966,0.0034147971,0.03966315],"category_scores_gemma":[0.010381321,0.0008612818,0.0026435417,0.0030283143,0.0006752862,0.0044157864,0.006061474,0.0038172274,0.07237917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056772697,0.00032584224,0.0014969308,0.0010848137,0.00013473346,0.00017123828,0.00020972751,0.0010467586,0.011863066,0.0007335465,0.9010417,0.08132389],"study_design_scores_gemma":[0.0011456923,0.0011340554,0.023430405,0.0005198138,0.0002987751,0.0014278498,0.0011418854,0.076844536,0.044858932,0.010874355,0.8379584,0.00036529687],"about_ca_topic_score_codex":0.0086919945,"about_ca_topic_score_gemma":0.013166261,"teacher_disagreement_score":0.03966315,"about_ca_system_score_codex":0.0018614411,"about_ca_system_score_gemma":0.0032243154,"threshold_uncertainty_score":0.1326865},"labels":[],"label_agreement":null},{"id":"W4389518395","doi":"10.18653/v1/2023.arabicnlp-1.83","title":"WojoodNER 2023: The First Arabic Named Entity Recognition Shared Task","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Arabic; Named-entity recognition; Computer science; Task (project management); Natural language processing; Focus (optics); Artificial intelligence; Linguistics; Engineering","score_opus":0.04759203722951118,"score_gpt":0.24432034228399283,"score_spread":0.19672830505448163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518395","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2692163,0.009052356,0.29234397,0.005553034,0.012101051,0.0091152685,0.16339384,0.1134384,0.12578575],"genre_scores_gemma":[0.20079945,0.00078425684,0.2763116,0.0023882047,0.0008849139,0.0049757618,0.44553682,0.008683167,0.059635937],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9917578,0.0026989703,0.0007411579,0.0023129093,0.0015406981,0.0009483479],"domain_scores_gemma":[0.98889625,0.0028088281,0.00033606467,0.0033893264,0.0025736275,0.0019958036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008066854,0.0041592526,0.0030107812,0.0029624428,0.0036374535,0.0036981637,0.0037111433,0.0044319555,0.018894024],"category_scores_gemma":[0.016939765,0.0008171962,0.002833738,0.0020767092,0.0013061111,0.007777441,0.013792913,0.0040942053,0.025502319],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025816984,0.0015892813,0.004183877,0.0023841707,0.0006742275,0.0015049725,0.0021024148,0.00633261,0.034787197,0.004512524,0.5126649,0.42668205],"study_design_scores_gemma":[0.0012047495,0.0027721594,0.020258825,0.00069595943,0.00048542587,0.0030315344,0.004297043,0.08438867,0.06373073,0.021782693,0.79651767,0.00083451596],"about_ca_topic_score_codex":0.011015524,"about_ca_topic_score_gemma":0.020443449,"teacher_disagreement_score":0.018894024,"about_ca_system_score_codex":0.0015426435,"about_ca_system_score_gemma":0.0046788747,"threshold_uncertainty_score":0.06320679},"labels":[],"label_agreement":null},{"id":"W4389518435","doi":"10.18653/v1/2023.banglalp-1.43","title":"Z-Index at BLP-2023 Task 2: A Comparative Study on Sentiment Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Lexical analysis; Computer science; Task (project management); Index (typography); Artificial intelligence; Natural language processing; Sentiment analysis; Class (philosophy); Machine learning; World Wide Web; Engineering","score_opus":0.06926382760000373,"score_gpt":0.3334456753192457,"score_spread":0.26418184771924197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518435","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87072146,0.007955247,0.040752057,0.0017242542,0.0020569433,0.0011468726,0.02795827,0.017426716,0.030258132],"genre_scores_gemma":[0.8437766,0.0011785625,0.039843917,0.00059663877,0.0007008404,0.0008468704,0.09768775,0.0023427575,0.013026135],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.994484,0.0021912067,0.00042596567,0.0011302612,0.0013088279,0.00045976858],"domain_scores_gemma":[0.9912538,0.0044712634,0.00039530944,0.0013228757,0.0018618801,0.0006950334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00618304,0.0024246962,0.0014777713,0.0027634138,0.0012586232,0.002463587,0.0014273253,0.0018151591,0.0060398793],"category_scores_gemma":[0.014034136,0.00032700424,0.0011872613,0.0020405108,0.0005176073,0.003956749,0.0025856404,0.0016209941,0.0075914217],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008258714,0.00391013,0.06263602,0.0038261362,0.001373346,0.0010402227,0.003013181,0.012686821,0.08193194,0.0013825686,0.20577478,0.61416626],"study_design_scores_gemma":[0.0014375958,0.009351524,0.35980743,0.0004889249,0.0010869466,0.0026259006,0.0074689784,0.3807714,0.08123435,0.006154813,0.14890847,0.0006636942],"about_ca_topic_score_codex":0.0058618584,"about_ca_topic_score_gemma":0.007402394,"teacher_disagreement_score":0.00618304,"about_ca_system_score_codex":0.0009767188,"about_ca_system_score_gemma":0.00088692986,"threshold_uncertainty_score":0.032699466},"labels":[],"label_agreement":null},{"id":"W4389518635","doi":"10.18653/v1/2023.findings-emnlp.776","title":"Detrimental Contexts in Open-Domain Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Question answering; Leverage (statistics); Pipeline (software); Open domain; Context (archaeology); Artificial intelligence; Information retrieval; Natural language processing; Matching (statistics); Machine learning; Language model","score_opus":0.03086295426828448,"score_gpt":0.3123762498647927,"score_spread":0.2815132955965082,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3926543,0.011283477,0.5820554,0.0018061195,0.00022105098,0.0002537096,0.00068078854,0.005552718,0.005492446],"genre_scores_gemma":[0.91789305,0.00096947374,0.07770699,0.00044757762,0.00017791103,0.00014418438,0.00083909964,0.00024698497,0.0015747567],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99669915,0.0018923946,0.00014380827,0.00079164834,0.00029691416,0.00017611089],"domain_scores_gemma":[0.9872801,0.01016333,0.0005480044,0.0010903483,0.00064051134,0.0002777403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045438744,0.0011385742,0.0009958365,0.001164956,0.0012247575,0.0017194434,0.0013693463,0.0019824451,0.0012909189],"category_scores_gemma":[0.022574898,0.00079628854,0.0008812207,0.0007603293,0.0012970329,0.0042475667,0.0027861705,0.0025274782,0.00080395374],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021369872,0.0006319404,0.0359892,0.0013956538,0.0004438486,0.0015252387,0.005734622,0.51052654,0.029016413,0.036529496,0.010832151,0.3652379],"study_design_scores_gemma":[0.00005397746,0.0002648331,0.0036864025,0.0001176514,0.00014416878,0.00036588326,0.00047602697,0.9254527,0.009163505,0.05365146,0.006564648,0.000058658734],"about_ca_topic_score_codex":0.0044124504,"about_ca_topic_score_gemma":0.008500522,"teacher_disagreement_score":0.0045438744,"about_ca_system_score_codex":0.0008135197,"about_ca_system_score_gemma":0.0009326276,"threshold_uncertainty_score":0.024030566},"labels":[],"label_agreement":null},{"id":"W4389518686","doi":"10.18653/v1/2023.emnlp-main.330","title":"Just Ask for Calibration: Strategies for Eliciting Calibrated Confidence Scores from Language Models Fine-Tuned with Human Feedback","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Office of Naval Research; Canadian Institute for Advanced Research; National Science Foundation","keywords":"Ask price; Calibration; Computer science; Tian; Artificial intelligence; Statistics; Art; Mathematics; Literature","score_opus":0.07316623696227058,"score_gpt":0.30503333613674294,"score_spread":0.23186709917447235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518686","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.080199525,0.0005946994,0.9032762,0.0012696071,0.0002008854,0.00060831016,0.0006834388,0.00881356,0.0043537924],"genre_scores_gemma":[0.6746309,0.00015115198,0.32057875,0.00065893005,0.00009000453,0.0009388414,0.00094419654,0.00071598147,0.001291277],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9822886,0.012088723,0.00084930845,0.0025890097,0.0017672131,0.00041717765],"domain_scores_gemma":[0.83502567,0.13327844,0.005753972,0.015527527,0.008819737,0.0015946546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021911804,0.002538387,0.0014894586,0.001791164,0.00073263637,0.0030268922,0.0035612362,0.003567116,0.0067491187],"category_scores_gemma":[0.22734286,0.0011688091,0.0007423374,0.0014544624,0.0010701275,0.0056147706,0.004765255,0.0036191603,0.0029982363],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037254095,0.0016563217,0.018941483,0.0008958505,0.00065174134,0.00046545526,0.004837494,0.06739479,0.033084776,0.013196675,0.023571577,0.83157843],"study_design_scores_gemma":[0.00074181764,0.0007997878,0.0062528164,0.00019530942,0.0001918367,0.0002519766,0.0013109441,0.87741023,0.03308703,0.07301553,0.006429397,0.0003133018],"about_ca_topic_score_codex":0.0012168324,"about_ca_topic_score_gemma":0.0022189685,"teacher_disagreement_score":0.021911804,"about_ca_system_score_codex":0.0009548627,"about_ca_system_score_gemma":0.0012269662,"threshold_uncertainty_score":0.11588204},"labels":[],"label_agreement":null},{"id":"W4389518693","doi":"10.18653/v1/2023.emnlp-main.403","title":"Understanding the Inner-workings of Language Models Through Representation Dissimilarity","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"Pacific Northwest National Laboratory; Laboratory Directed Research and Development; Battelle; U.S. Department of Energy","keywords":"Interpretability; Computer science; Representation (politics); Language model; Generalization; Measure (data warehouse); Set (abstract data type); Transparency (behavior); Natural language processing; Artificial intelligence; Machine learning; Data mining; Mathematics; Programming language","score_opus":0.2546609340372178,"score_gpt":0.34455411413175957,"score_spread":0.08989318009454178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518693","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28291878,0.00036826788,0.7093296,0.0014979024,0.0000348909,0.000062773026,0.00013367043,0.0003418945,0.005312146],"genre_scores_gemma":[0.92846143,0.00014871334,0.07052887,0.000082728904,0.000027133174,0.0000609038,0.00012518452,0.00014471322,0.000420336],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9950294,0.0030497457,0.00027134895,0.0007310906,0.0006984283,0.00022002669],"domain_scores_gemma":[0.95482385,0.03236895,0.0034292697,0.0069205137,0.0014330567,0.0010243001],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010583857,0.0007562826,0.001140432,0.0020982437,0.001135178,0.0062086545,0.0015814713,0.0014824322,0.0024051024],"category_scores_gemma":[0.07021758,0.0008560536,0.001182214,0.0010079386,0.0051132995,0.01792755,0.0055469414,0.003525001,0.00030776067],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000700598,0.00022337407,0.039785776,0.00040118527,0.00042943645,0.0003709832,0.012632716,0.18517025,0.019346729,0.59710866,0.0011370592,0.14269319],"study_design_scores_gemma":[0.000022346689,0.00013661046,0.0037935646,0.000041596584,0.00004442297,0.00020484808,0.00092018757,0.5601246,0.0034174267,0.4297125,0.001519022,0.00006280609],"about_ca_topic_score_codex":0.0016061311,"about_ca_topic_score_gemma":0.0011436751,"teacher_disagreement_score":0.010583857,"about_ca_system_score_codex":0.0014374857,"about_ca_system_score_gemma":0.00073694595,"threshold_uncertainty_score":0.05597347},"labels":[],"label_agreement":null},{"id":"W4389518718","doi":"10.18653/v1/2023.emnlp-main.427","title":"Rather a Nurse than a Physician - Contrastive Explanations under Investigation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Berlin Center for Machine Learning; Novo Nordisk Fonden; Research Executive Agency; Banting and Best Diabetes Centre, University of Toronto; Novo Nordisk; European Commission","keywords":"Contrastive analysis; Computer science; Linguistics; Contrast (vision); Natural language processing; Artificial intelligence; Psychology","score_opus":0.038741324207895396,"score_gpt":0.2665800505663476,"score_spread":0.2278387263584522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518718","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26838854,0.009413654,0.5112266,0.045370776,0.0027424542,0.000732348,0.04082983,0.009356005,0.11193976],"genre_scores_gemma":[0.835893,0.0018710891,0.121444404,0.0053058146,0.0003999682,0.00024070464,0.013995073,0.00093526987,0.01991476],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9979493,0.00088771665,0.00013791688,0.0005291939,0.00039402518,0.00010190993],"domain_scores_gemma":[0.98610145,0.010264501,0.0011186899,0.0013260177,0.00089828495,0.00029096703],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032132145,0.00062937965,0.00031247636,0.0009921957,0.000820821,0.0015192927,0.000848776,0.00153977,0.017197538],"category_scores_gemma":[0.020686999,0.00031701635,0.0007262646,0.0010854267,0.00073595357,0.0027918161,0.00089906424,0.0018602632,0.0033630242],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018069003,0.00035821777,0.19135149,0.00495637,0.00050912885,0.0042832657,0.0128582,0.0106378365,0.02500107,0.14163268,0.1946278,0.41197702],"study_design_scores_gemma":[0.00022115966,0.00028696522,0.05454872,0.0012825711,0.00039290206,0.0051362175,0.0060897637,0.028940083,0.016691966,0.08988846,0.7963412,0.00017995763],"about_ca_topic_score_codex":0.0031712973,"about_ca_topic_score_gemma":0.00961089,"teacher_disagreement_score":0.017197538,"about_ca_system_score_codex":0.001009011,"about_ca_system_score_gemma":0.0013713007,"threshold_uncertainty_score":0.057531536},"labels":[],"label_agreement":null},{"id":"W4389518728","doi":"10.18653/v1/2023.emnlp-main.357","title":"RainProof: An Umbrella to Shield Text Generator from Out-Of-Distribution Data","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Shield; Computer science; Generator (circuit theory); Geology; Physics","score_opus":0.13003355009157463,"score_gpt":0.32678511806055344,"score_spread":0.1967515679689788,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518728","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005695877,0.0001745335,0.57711643,0.0007950269,0.0006470152,0.0005247615,0.008540679,0.39799076,0.008514956],"genre_scores_gemma":[0.19381271,0.00064877677,0.44177726,0.0030095286,0.0014305004,0.0021831472,0.038845792,0.24210997,0.07618235],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99688894,0.0010842512,0.00037129587,0.00055330736,0.0008747312,0.00022754555],"domain_scores_gemma":[0.97716796,0.010212454,0.0009852907,0.007904506,0.0026057446,0.0011240928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054016686,0.0022863562,0.0012715262,0.0022222204,0.00090461824,0.0036202557,0.0026801073,0.0019783077,0.06945919],"category_scores_gemma":[0.027485486,0.0016354644,0.0011575331,0.001408662,0.00090033386,0.004858033,0.0050406866,0.0024078086,0.048133317],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004089757,0.0004085652,0.0047778455,0.001020125,0.00028152135,0.0012396002,0.0019625507,0.0065700235,0.052602112,0.013569847,0.53781873,0.37565932],"study_design_scores_gemma":[0.0008891784,0.00058560877,0.002703517,0.00028682765,0.0002594531,0.0013354787,0.0003452508,0.15169194,0.1974966,0.028116629,0.6160314,0.00025818244],"about_ca_topic_score_codex":0.0007441672,"about_ca_topic_score_gemma":0.00070127496,"teacher_disagreement_score":0.06945919,"about_ca_system_score_codex":0.0005942097,"about_ca_system_score_gemma":0.0011555612,"threshold_uncertainty_score":0.23236418},"labels":[],"label_agreement":null},{"id":"W4389518736","doi":"10.18653/v1/2023.findings-emnlp.975","title":"Representation Projection Invariance Mitigates Representation Collapse","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Allen Institute for Artificial Intelligence; University of Toronto; Johns Hopkins University","keywords":"Computer science; Regularization (linguistics); Representation (politics); Robustness (evolution); Benchmark (surveying); Artificial intelligence; Machine learning; Algorithm; Theoretical computer science","score_opus":0.06431351550441888,"score_gpt":0.3202886456303431,"score_spread":0.2559751301259242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518736","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.116497174,0.0010412249,0.8711305,0.00058925844,0.000118874144,0.00017974984,0.00029978566,0.007474656,0.002668808],"genre_scores_gemma":[0.72831976,0.0004328367,0.26355657,0.00069133274,0.00014809,0.00029471415,0.0017569284,0.0010465676,0.003753243],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99726295,0.0009479115,0.00014584837,0.00085018267,0.0005411075,0.0002519438],"domain_scores_gemma":[0.9946239,0.0022065481,0.00045002613,0.0019644762,0.00055478304,0.00020031683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032257345,0.0018031141,0.0012528719,0.0009927469,0.0008542581,0.0014186915,0.0021888523,0.00157857,0.001814386],"category_scores_gemma":[0.014047708,0.0005210154,0.0011188937,0.0009515157,0.0014302881,0.0033648196,0.0040575857,0.004017615,0.0010806003],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053104764,0.00075051776,0.0052914377,0.00044308172,0.00028815633,0.00031310978,0.0008048026,0.23661815,0.083319314,0.01430252,0.014828851,0.642509],"study_design_scores_gemma":[0.000046602083,0.00026884637,0.0016167305,0.000034547807,0.000070242,0.00021419587,0.00016248113,0.95267755,0.024503171,0.015756007,0.004605887,0.00004372403],"about_ca_topic_score_codex":0.0035052453,"about_ca_topic_score_gemma":0.0049296,"teacher_disagreement_score":0.0035052453,"about_ca_system_score_codex":0.0008681158,"about_ca_system_score_gemma":0.0015351006,"threshold_uncertainty_score":0.017059505},"labels":[],"label_agreement":null},{"id":"W4389518777","doi":"10.18653/v1/2023.emnlp-main.257","title":"Transductive Learning for Textual Few-Shot Classification in API-based Embedding Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Grand Équipement National De Calcul Intensif","keywords":"Computer science; Embedding; Artificial intelligence; Natural language processing; Shot (pellet); Machine learning; Chemistry","score_opus":0.1290711266498333,"score_gpt":0.33144356573395756,"score_spread":0.20237243908412425,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518777","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06309764,0.003862728,0.9241649,0.0015016856,0.0003010563,0.00022642521,0.0010147287,0.0028168894,0.0030137976],"genre_scores_gemma":[0.85319865,0.001034324,0.12702297,0.0008274394,0.0007976911,0.00067409384,0.0047151586,0.000667127,0.01106251],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99772626,0.0010696441,0.00013431686,0.0006155347,0.00024678346,0.00020746005],"domain_scores_gemma":[0.98853266,0.008971511,0.00040067043,0.0009772776,0.0008075569,0.00031036735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038188219,0.0014291253,0.002535513,0.0018665129,0.00091904163,0.0023857448,0.004016658,0.0026076487,0.003873628],"category_scores_gemma":[0.013140549,0.0008309227,0.001350738,0.0017394967,0.0012246416,0.0060650054,0.0027252224,0.0040908563,0.0024119618],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008826593,0.0011675619,0.004422606,0.00062908634,0.00035972163,0.00022021787,0.00065107027,0.37357703,0.0026537203,0.032635834,0.026466329,0.5563342],"study_design_scores_gemma":[0.00001046838,0.000034813635,0.00017146561,0.000015435318,0.000012947069,0.000016509677,0.000027488473,0.9770098,0.0002625978,0.021995628,0.0004329417,0.000009927999],"about_ca_topic_score_codex":0.0043372805,"about_ca_topic_score_gemma":0.0066441647,"teacher_disagreement_score":0.0043372805,"about_ca_system_score_codex":0.0017207724,"about_ca_system_score_gemma":0.00080067205,"threshold_uncertainty_score":0.02019614},"labels":[],"label_agreement":null},{"id":"W4389518858","doi":"10.18653/v1/2023.findings-emnlp.674","title":"Balaur: Language Model Pretraining with Lexical Semantic Relations","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Samsung; Nvidia","keywords":"Computer science; Natural language processing; Inference; Artificial intelligence; Generalization; Transformer; Set (abstract data type); Language model; Meaning (existential); Interface (matter); Programming language; Psychology","score_opus":0.029940624321397952,"score_gpt":0.26493709039858887,"score_spread":0.2349964660771909,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518858","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027925711,0.00075811177,0.93038964,0.000668823,0.00031623873,0.00015130341,0.0011208659,0.034288343,0.0043810396],"genre_scores_gemma":[0.4706205,0.00045089138,0.50250953,0.0013540622,0.00022741918,0.00065908395,0.006728627,0.003748091,0.0137016345],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99949753,0.00015693452,0.000026160422,0.00020506006,0.00006160526,0.00005263988],"domain_scores_gemma":[0.99853325,0.0009753905,0.000047010366,0.00023771144,0.00014114311,0.00006548077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011611212,0.001458518,0.0008740108,0.0006859148,0.0004803435,0.0011658237,0.0021793947,0.0013630936,0.011089826],"category_scores_gemma":[0.004659785,0.0007644216,0.0012640385,0.00060145144,0.00061612803,0.0028866795,0.002109223,0.0044244844,0.005295587],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006227451,0.0003361717,0.0022553958,0.0004784427,0.0003567435,0.00026259772,0.00053938676,0.19596799,0.023820966,0.014290804,0.044401743,0.716667],"study_design_scores_gemma":[0.00004635454,0.000060507442,0.00033096675,0.000028090124,0.000034858982,0.00005614332,0.00006939706,0.9758523,0.005692228,0.013455359,0.004350763,0.00002295467],"about_ca_topic_score_codex":0.005281685,"about_ca_topic_score_gemma":0.014419465,"teacher_disagreement_score":0.011089826,"about_ca_system_score_codex":0.00069986423,"about_ca_system_score_gemma":0.00110817,"threshold_uncertainty_score":0.037099183},"labels":[],"label_agreement":null},{"id":"W4389518861","doi":"10.18653/v1/2023.findings-emnlp.616","title":"Knowledge Corpus Error in Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Pipeline (software); Computer science; Context (archaeology); Natural language processing; String (physics); Artificial intelligence; Space (punctuation); Question answering; Domain (mathematical analysis); Information retrieval; History; Mathematics","score_opus":0.051404789048701154,"score_gpt":0.31377550380935676,"score_spread":0.2623707147606556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518861","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33651438,0.012414808,0.61577815,0.0066862926,0.001042375,0.00063934224,0.0019025081,0.010358174,0.014664084],"genre_scores_gemma":[0.85130996,0.0012358649,0.13725711,0.0015853127,0.00032722234,0.00034917748,0.0032237836,0.0012435135,0.0034681205],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96684444,0.01973366,0.0019559579,0.005270053,0.005541323,0.0006546209],"domain_scores_gemma":[0.8003322,0.16853414,0.005271078,0.015284934,0.009794607,0.0007831314],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019147852,0.00082207844,0.0013675446,0.0019756781,0.0017065014,0.002876962,0.0018333407,0.0026896326,0.0023323426],"category_scores_gemma":[0.15609184,0.0006889832,0.00069557043,0.0025532928,0.0026376871,0.00606366,0.004434144,0.0024975354,0.0011364871],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020478081,0.0006624845,0.03890014,0.0042118793,0.00053803856,0.0021163365,0.014977155,0.10355838,0.028594818,0.066124275,0.041550208,0.69671845],"study_design_scores_gemma":[0.0003044251,0.0010419552,0.019076673,0.0014924444,0.00060899276,0.004743411,0.0046883696,0.57240725,0.115561634,0.16290794,0.11685773,0.00030907762],"about_ca_topic_score_codex":0.0052281423,"about_ca_topic_score_gemma":0.0032934851,"teacher_disagreement_score":0.019147852,"about_ca_system_score_codex":0.0017893091,"about_ca_system_score_gemma":0.0018407746,"threshold_uncertainty_score":0.101264775},"labels":[],"label_agreement":null},{"id":"W4389518871","doi":"10.18653/v1/2023.findings-emnlp.579","title":"COUNT: COntrastive UNlikelihood Text Style Transfer for Text Detoxification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Fluency; Artificial intelligence; Natural language processing; Focus (optics); Classifier (UML); Machine learning; Mathematics; Mathematics education","score_opus":0.028055416724721623,"score_gpt":0.2601129224531013,"score_spread":0.2320575057283797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518871","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33763236,0.004922978,0.5957657,0.0016654459,0.0012144793,0.00090286025,0.0054917564,0.036145058,0.01625933],"genre_scores_gemma":[0.66589046,0.00079527893,0.2844787,0.0013385158,0.00064583437,0.0009382676,0.015840163,0.0026274936,0.02744525],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982528,0.00060011394,0.000109627086,0.0005623625,0.00036489323,0.00011028635],"domain_scores_gemma":[0.9962508,0.0018477002,0.00030464312,0.0008482234,0.0005189842,0.00022963762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030783813,0.002175141,0.0009500816,0.0015616042,0.0007403684,0.0015956158,0.0020424945,0.0019594466,0.0048081577],"category_scores_gemma":[0.008771444,0.00028352984,0.0012030362,0.00089446903,0.00113015,0.0033303734,0.0026787415,0.0022708846,0.0037016687],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015354687,0.0009801971,0.007191678,0.0009802732,0.00032252743,0.00038769396,0.0006170496,0.05498644,0.062576756,0.0069081103,0.04062441,0.82288945],"study_design_scores_gemma":[0.00036093703,0.0016296201,0.0067935777,0.00011915898,0.00015078558,0.0008594417,0.00044252252,0.8632324,0.07847704,0.019484846,0.02830092,0.0001487571],"about_ca_topic_score_codex":0.0012875937,"about_ca_topic_score_gemma":0.003650797,"teacher_disagreement_score":0.0048081577,"about_ca_system_score_codex":0.00090031716,"about_ca_system_score_gemma":0.0008231449,"threshold_uncertainty_score":0.016280234},"labels":[],"label_agreement":null},{"id":"W4389518890","doi":"10.18653/v1/2023.findings-emnlp.634","title":"Mixture-of-Linguistic-Experts Adapters for Improving and Interpreting Pre-trained Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of the Fraser Valley; University of British Columbia","funders":"","keywords":"Softmax function; Computer science; Language model; Gumbel distribution; Adapter (computing); Layer (electronics); Artificial intelligence; Pruning; Natural language processing; Encoding (memory); Artificial neural network; Mathematics; Statistics; Extreme value theory","score_opus":0.020131400302049834,"score_gpt":0.2764942223730839,"score_spread":0.25636282207103406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518890","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015221728,0.00039431977,0.9773933,0.00017288157,0.000062930354,0.00007706029,0.00008490783,0.0052718665,0.0013209312],"genre_scores_gemma":[0.3756603,0.0005081143,0.6117528,0.00061931333,0.00018858338,0.00038332882,0.0011236515,0.0010778651,0.00868614],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986634,0.00043779705,0.00007617758,0.00044008013,0.00023834537,0.00014411243],"domain_scores_gemma":[0.9973236,0.0013303187,0.00016903624,0.00055961765,0.0004795377,0.00013795533],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033768548,0.0025680757,0.0011851173,0.0017270928,0.0006285995,0.0016238318,0.003299846,0.0021882234,0.005494905],"category_scores_gemma":[0.009425808,0.0012088381,0.001818099,0.0011846506,0.0010336239,0.004487673,0.0034054127,0.004232154,0.0032936493],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064555305,0.00025830546,0.0032342777,0.00024032246,0.00041188675,0.00031136096,0.00079869747,0.21853833,0.02478127,0.014545295,0.006855562,0.7293791],"study_design_scores_gemma":[0.000022146149,0.000060890037,0.00039457704,0.000023952365,0.00007318716,0.0000887423,0.000055618155,0.97714907,0.009653242,0.0096168,0.002836193,0.000025604133],"about_ca_topic_score_codex":0.0054946057,"about_ca_topic_score_gemma":0.011640223,"teacher_disagreement_score":0.005494905,"about_ca_system_score_codex":0.0011578801,"about_ca_system_score_gemma":0.0013884366,"threshold_uncertainty_score":0.018382251},"labels":[],"label_agreement":null},{"id":"W4389518966","doi":"10.18653/v1/2023.findings-emnlp.285","title":"Measuring the Knowledge Acquisition-Utilization Gap in Pretrained Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute; Canadian Institute for Advanced Research; McGill University","funders":"","keywords":"Computer science; USable; Domain knowledge; Robustness (evolution); Parametric statistics; Task (project management); Measure (data warehouse); Knowledge acquisition; Artificial intelligence; Knowledge extraction; Machine learning; Data mining; Engineering; Mathematics","score_opus":0.12079531531899312,"score_gpt":0.30065930140701147,"score_spread":0.17986398608801835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518966","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9420607,0.00090911816,0.051405907,0.00035382804,0.000028884275,0.000092503135,0.0007420477,0.001285728,0.003121368],"genre_scores_gemma":[0.9813844,0.00017840668,0.016347686,0.00008072292,0.000014012846,0.00008185638,0.0013984278,0.000098513876,0.00041593375],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9954112,0.0018319947,0.00046579976,0.0012126964,0.00072244246,0.00035598865],"domain_scores_gemma":[0.9324485,0.05188706,0.0034123468,0.00886152,0.0020542387,0.001336256],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010178699,0.0012473435,0.0010516414,0.0015972614,0.00058428943,0.0025355755,0.0013712621,0.0016791878,0.001631437],"category_scores_gemma":[0.06879868,0.00067849626,0.000848899,0.0011205962,0.0010829134,0.007770457,0.0037302333,0.0030715624,0.0006950214],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033934414,0.0020386032,0.22938861,0.0011707282,0.0016269197,0.0005319541,0.004237782,0.25533068,0.038664263,0.0052804244,0.0034899197,0.45484683],"study_design_scores_gemma":[0.00012420032,0.0029161954,0.122560255,0.00020489791,0.0005570627,0.00062882446,0.0018327499,0.79812264,0.049788345,0.01950756,0.0035176857,0.00023962007],"about_ca_topic_score_codex":0.0034835006,"about_ca_topic_score_gemma":0.0038836405,"teacher_disagreement_score":0.010178699,"about_ca_system_score_codex":0.00095591776,"about_ca_system_score_gemma":0.0014876059,"threshold_uncertainty_score":0.053830743},"labels":[],"label_agreement":null},{"id":"W4389518973","doi":"10.18653/v1/2023.findings-emnlp.315","title":"Improving Contrastive Learning of Sentence Embeddings with Focal InfoNCE","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Sentence; Computer science; Exploit; Representation (politics); Artificial intelligence; Natural language processing; Function (biology); Quality (philosophy)","score_opus":0.011613599491793091,"score_gpt":0.23067607182902436,"score_spread":0.21906247233723128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518973","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14153656,0.0009771448,0.8519019,0.0005151107,0.00014353685,0.00009901118,0.00046887764,0.001981548,0.0023763874],"genre_scores_gemma":[0.7995595,0.00033507426,0.19094133,0.00053369545,0.00024327911,0.00018837469,0.002968115,0.00041224904,0.0048184963],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99869484,0.0005435728,0.00007784891,0.00036900412,0.00022997102,0.00008477355],"domain_scores_gemma":[0.9952236,0.0029047767,0.0003221757,0.00060275424,0.000799527,0.00014712545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022361404,0.0012391661,0.0010458805,0.0012591496,0.0004269109,0.0010454525,0.0011922874,0.0010790357,0.0021081965],"category_scores_gemma":[0.011046977,0.00035214744,0.0007585779,0.0010246577,0.00088776986,0.002915129,0.0020368544,0.0022111372,0.0010755355],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088668114,0.000773557,0.010283234,0.0004655801,0.00029923904,0.00031730582,0.0007070006,0.15072593,0.043403782,0.024977824,0.0129221715,0.7542377],"study_design_scores_gemma":[0.000033041655,0.00028976105,0.0012259997,0.000025381967,0.000040831048,0.000110879104,0.000082479964,0.9701876,0.010118753,0.016084064,0.0017790714,0.00002217618],"about_ca_topic_score_codex":0.0009746207,"about_ca_topic_score_gemma":0.0023458581,"teacher_disagreement_score":0.0022361404,"about_ca_system_score_codex":0.00063190947,"about_ca_system_score_gemma":0.00073709205,"threshold_uncertainty_score":0.011825979},"labels":[],"label_agreement":null},{"id":"W4389518978","doi":"10.18653/v1/2023.findings-emnlp.393","title":"Aligning Language Models to User Opinions","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Persona; Ideology; Demographics; Public opinion; Set (abstract data type); Computer science; User group; Work (physics); Data science; Internet privacy; Human–computer interaction; World Wide Web; Political science; Sociology; Engineering; Politics","score_opus":0.06393325834733068,"score_gpt":0.30317322365692123,"score_spread":0.23923996530959055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518978","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16377543,0.00080750504,0.81943893,0.0018971014,0.0002507626,0.0002556626,0.0014291601,0.0069997204,0.005145729],"genre_scores_gemma":[0.79406,0.00031605,0.19733737,0.0005737503,0.00018053554,0.00030787487,0.0033352368,0.00072888297,0.0031601714],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9948102,0.003197661,0.0002711963,0.0009160635,0.0005963245,0.0002085813],"domain_scores_gemma":[0.9896576,0.006408096,0.00064542197,0.0012378512,0.0017694052,0.00028161786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051289555,0.0009985408,0.00075567485,0.0019468562,0.0004824304,0.0028856532,0.0012099423,0.0012437109,0.002893964],"category_scores_gemma":[0.028299753,0.00056192954,0.0010694982,0.0012859373,0.0004653583,0.0036343269,0.0019333177,0.002095505,0.0028700598],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011282688,0.00068448076,0.039998583,0.00061452086,0.00062140595,0.00040494124,0.0037471647,0.2470621,0.0290805,0.017465094,0.017325012,0.64186794],"study_design_scores_gemma":[0.000018478015,0.00010845404,0.002143463,0.00003916932,0.000058467635,0.00005420404,0.000493279,0.969454,0.0045635044,0.016928373,0.0061006104,0.000038005655],"about_ca_topic_score_codex":0.0040717307,"about_ca_topic_score_gemma":0.005522279,"teacher_disagreement_score":0.0051289555,"about_ca_system_score_codex":0.000981356,"about_ca_system_score_gemma":0.0009682213,"threshold_uncertainty_score":0.027124822},"labels":[],"label_agreement":null},{"id":"W4389518993","doi":"10.18653/v1/2023.findings-emnlp.511","title":"Evaluating Dependencies in Fact Editing for Language Models: Specificity and Implication Awareness","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research; Samsung; Nvidia","keywords":"Dependency (UML); Computer science; Knowledge base; Protocol (science); Natural language processing; Neglect; Process (computing); Artificial intelligence; Data science; Programming language; Psychology","score_opus":0.18229258343695362,"score_gpt":0.390164939008229,"score_spread":0.2078723555712754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389518993","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6585734,0.0021866034,0.3165986,0.0013045298,0.0001551292,0.000877654,0.0049653156,0.0090419,0.0062969583],"genre_scores_gemma":[0.8123714,0.00025546664,0.17886269,0.00024098542,0.00006266441,0.00023076142,0.006923364,0.00036922528,0.0006834474],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98636925,0.006954496,0.0009494041,0.0027143005,0.0026087833,0.0004038886],"domain_scores_gemma":[0.78565705,0.1844533,0.0070240507,0.017061466,0.004297157,0.0015069091],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018268434,0.001025081,0.00074061885,0.002836921,0.0009627219,0.0030893555,0.0021427404,0.0022214127,0.001363376],"category_scores_gemma":[0.12246953,0.000680413,0.0012165519,0.0015874956,0.0012667599,0.007706984,0.0035069694,0.0031334579,0.00036815947],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031110742,0.001979137,0.13559802,0.0027534836,0.0013568302,0.0016316036,0.008901158,0.16214578,0.03350116,0.021651594,0.018821482,0.6085487],"study_design_scores_gemma":[0.00020607801,0.0006503562,0.029025778,0.0001983563,0.00049205637,0.0010313182,0.0017301416,0.8673914,0.035295248,0.04881203,0.015017193,0.000150069],"about_ca_topic_score_codex":0.0053457655,"about_ca_topic_score_gemma":0.010377324,"teacher_disagreement_score":0.018268434,"about_ca_system_score_codex":0.0013532083,"about_ca_system_score_gemma":0.0017982065,"threshold_uncertainty_score":0.096613884},"labels":[],"label_agreement":null},{"id":"W4389519056","doi":"10.18653/v1/2023.emnlp-main.435","title":"A Mechanistic Interpretation of Arithmetic Reasoning in Language Models using Causal Mediation Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; Open Philanthropy Project; National Science Foundation","keywords":"Computer science; Mediation; Security token; Causal reasoning; Set (abstract data type); Interpretation (philosophy); Theoretical computer science; Natural language processing; Process (computing); Language model; Causal model; Artificial intelligence; Arithmetic; Cognition; Programming language; Mathematics; Psychology","score_opus":0.02509132893953309,"score_gpt":0.29085532487544863,"score_spread":0.26576399593591554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519056","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07136711,0.00013459795,0.9171103,0.00172046,0.000051514733,0.00010377403,0.00019326493,0.0009502796,0.008368628],"genre_scores_gemma":[0.90960634,0.00012818039,0.087954335,0.00020844176,0.000036386446,0.00013960275,0.00015709033,0.000103797094,0.0016657113],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990589,0.00047058804,0.000044530676,0.00018697137,0.0001413248,0.00009763233],"domain_scores_gemma":[0.99646336,0.002344323,0.00034326626,0.0004712,0.00023253597,0.000145231],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020256282,0.0006397271,0.000535,0.00084476155,0.0005240753,0.002281077,0.0019523212,0.0010506465,0.008795614],"category_scores_gemma":[0.01085735,0.00066353066,0.0013977288,0.00046845758,0.0020818817,0.004845805,0.0021859258,0.0022286198,0.00055167614],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003811761,0.00020457215,0.004385484,0.00028859853,0.00027557954,0.00070679246,0.0023231201,0.18695444,0.03133839,0.7162112,0.001978823,0.054951843],"study_design_scores_gemma":[0.000032183765,0.00006736746,0.000815667,0.000017896315,0.00005396098,0.000106671185,0.00013396687,0.5633747,0.0037789866,0.4307334,0.0008526081,0.0000325728],"about_ca_topic_score_codex":0.0026091943,"about_ca_topic_score_gemma":0.0018027604,"teacher_disagreement_score":0.008795614,"about_ca_system_score_codex":0.0010721664,"about_ca_system_score_gemma":0.0009928179,"threshold_uncertainty_score":0.02942425},"labels":[],"label_agreement":null},{"id":"W4389519096","doi":"10.18653/v1/2023.emnlp-main.439","title":"Prompt-Based Monte-Carlo Tree Search for Goal-oriented Dialogue Policy Planning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Monte Carlo tree search; Computer science; Task (project management); Tree (set theory); Machine learning; Artificial intelligence; Resource (disambiguation); Monte Carlo method; Engineering","score_opus":0.058478333879909646,"score_gpt":0.32376339571412255,"score_spread":0.2652850618342129,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519096","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048510846,0.00043607297,0.9380215,0.00046573865,0.00011931182,0.00017654951,0.00026287744,0.007611136,0.0043960046],"genre_scores_gemma":[0.7392115,0.00011964122,0.25686985,0.0002829857,0.00003890707,0.00032950434,0.00055262464,0.0005168513,0.0020781073],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988978,0.0005897193,0.00005259615,0.00022313822,0.00015872487,0.00007800364],"domain_scores_gemma":[0.9933031,0.005556711,0.00018228633,0.0003270005,0.00040658007,0.00022427249],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022410026,0.00085533765,0.0008566447,0.00058137585,0.00047480018,0.00091864815,0.0015363488,0.0014712512,0.00770278],"category_scores_gemma":[0.015276428,0.00056096766,0.0004748348,0.00048661986,0.00071315584,0.0018273845,0.0013578495,0.0021137584,0.0012595891],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009893903,0.00037540344,0.003153474,0.0004077141,0.00008414762,0.00017483998,0.00063628337,0.79254717,0.005045313,0.018890195,0.0072360015,0.17046],"study_design_scores_gemma":[0.000032061154,0.00003057284,0.00009600224,0.000007657226,0.000005490899,0.000009334706,0.0000191645,0.9940883,0.00059491163,0.004568889,0.0005425599,0.0000051447123],"about_ca_topic_score_codex":0.006269207,"about_ca_topic_score_gemma":0.0105850315,"teacher_disagreement_score":0.00770278,"about_ca_system_score_codex":0.0010803898,"about_ca_system_score_gemma":0.0022401945,"threshold_uncertainty_score":0.02576834},"labels":[],"label_agreement":null},{"id":"W4389519117","doi":"10.18653/v1/2023.emnlp-main.491","title":"Don’t Trust ChatGPT when your Question is not in English: A Study of Multilingual Abilities and Types of LLMs","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Machine Intelligence Institute","keywords":"Variety (cybernetics); Computer science; Natural language processing; Artificial intelligence; Linguistics","score_opus":0.04011867324513748,"score_gpt":0.2954704506340845,"score_spread":0.25535177738894704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519117","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9919743,0.00015993588,0.0037922498,0.00027189773,0.000016449232,0.0000597944,0.00014841999,0.00015935385,0.0034175813],"genre_scores_gemma":[0.996748,0.000082096485,0.002194165,0.00013548168,0.000015337328,0.000040201798,0.00020519651,0.00006749692,0.00051204767],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9920844,0.0047531496,0.00061682414,0.0012961287,0.0008848685,0.00036456544],"domain_scores_gemma":[0.84637135,0.12758544,0.008513856,0.010972429,0.0042547057,0.0023022196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009837916,0.0005269083,0.0006321197,0.0012137329,0.00077477243,0.0021489481,0.0010938462,0.0011166678,0.0021320954],"category_scores_gemma":[0.1022288,0.00047493484,0.00047154396,0.0009596175,0.0019976262,0.0066151735,0.0029806078,0.0020097347,0.0008427054],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032631804,0.0016163107,0.58090633,0.001391296,0.00061181304,0.0028386456,0.10223966,0.010358668,0.024886373,0.0070199296,0.0051265433,0.2597413],"study_design_scores_gemma":[0.00030417016,0.003658043,0.64392865,0.0005380863,0.0007335966,0.005778635,0.06817988,0.19186904,0.030851247,0.0315331,0.021999637,0.0006259478],"about_ca_topic_score_codex":0.0034897863,"about_ca_topic_score_gemma":0.0034280368,"teacher_disagreement_score":0.009837916,"about_ca_system_score_codex":0.00071631774,"about_ca_system_score_gemma":0.00077275874,"threshold_uncertainty_score":0.052028477},"labels":[],"label_agreement":null},{"id":"W4389519126","doi":"10.18653/v1/2023.emnlp-main.448","title":"OssCSE: Overcoming Surface Structure Bias in Contrastive Learning for Unsupervised Sentence Embedding","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Natural language processing; Embedding; Sentence; Computer science; Artificial intelligence; Linguistics; Philosophy","score_opus":0.05164863126802949,"score_gpt":0.30093111289566843,"score_spread":0.24928248162763894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519126","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024348393,0.00071612175,0.9586677,0.00030853046,0.00024982842,0.00016949371,0.0008987148,0.011860581,0.00278063],"genre_scores_gemma":[0.27554327,0.00046287017,0.7020111,0.00048906694,0.00028322087,0.0005519893,0.007438653,0.0027131233,0.010506787],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99892896,0.00042690255,0.000062156716,0.0003070307,0.00019731709,0.00007765433],"domain_scores_gemma":[0.9974579,0.0012737375,0.00009066924,0.0006889905,0.00037756696,0.000111152054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018976616,0.0015431532,0.0013103536,0.0012611216,0.00082103454,0.0012832095,0.0023501338,0.0013002761,0.008199535],"category_scores_gemma":[0.0063383747,0.00065668003,0.0010328837,0.0012381409,0.00094866374,0.005000766,0.0035838955,0.0025601694,0.003422891],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046497604,0.00028390661,0.0017811744,0.00032825096,0.00020321205,0.00015720914,0.00049049535,0.03195637,0.019781822,0.027542682,0.041558273,0.8754516],"study_design_scores_gemma":[0.0001414422,0.00017507891,0.00060626806,0.00003001137,0.000054884826,0.00010318854,0.00013548831,0.9281309,0.011647864,0.0475901,0.011348598,0.0000360434],"about_ca_topic_score_codex":0.0037253313,"about_ca_topic_score_gemma":0.012199792,"teacher_disagreement_score":0.008199535,"about_ca_system_score_codex":0.0006593302,"about_ca_system_score_gemma":0.0011552244,"threshold_uncertainty_score":0.027430177},"labels":[],"label_agreement":null},{"id":"W4389519204","doi":"10.18653/v1/2023.emnlp-main.721","title":"Reward-Augmented Decoding: Efficient Controlled Text Generation With a Unidirectional Reward Model","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Language model; Decoding methods; Overhead (engineering); Cache; Artificial intelligence; Text generation; Algorithm; Parallel computing; Programming language","score_opus":0.04398754862890349,"score_gpt":0.256004301114469,"score_spread":0.21201675248556548,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519204","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025777875,0.00026212892,0.9609732,0.00035803329,0.00012873264,0.00016721604,0.00024544337,0.008841256,0.0032460166],"genre_scores_gemma":[0.52076983,0.00019319348,0.46760336,0.00042330468,0.00012028109,0.0004641756,0.0009460065,0.0017976274,0.007682318],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99888307,0.00049380667,0.00006935284,0.0002444372,0.00022794203,0.000081512604],"domain_scores_gemma":[0.99534154,0.0033147018,0.00018933779,0.0005454158,0.00048712816,0.00012189688],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015748574,0.0010091917,0.0008134197,0.0005110598,0.00041218117,0.000934768,0.0014841789,0.0011828295,0.0042136987],"category_scores_gemma":[0.011573086,0.00040367217,0.0005386141,0.00048629552,0.0006875087,0.0015812275,0.0013945777,0.0015308828,0.0026583134],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058442494,0.00024322898,0.002356676,0.00029833225,0.00006623908,0.00043127532,0.0005506726,0.4209057,0.02929158,0.02554112,0.013417634,0.50631315],"study_design_scores_gemma":[0.00005022593,0.000047820504,0.00008832715,0.0000094855195,0.000011150887,0.00005122779,0.000016168127,0.9820961,0.0064208717,0.009274275,0.0019201979,0.0000142108265],"about_ca_topic_score_codex":0.0022300165,"about_ca_topic_score_gemma":0.003746161,"teacher_disagreement_score":0.0042136987,"about_ca_system_score_codex":0.0005809521,"about_ca_system_score_gemma":0.0012009443,"threshold_uncertainty_score":0.0140962005},"labels":[],"label_agreement":null},{"id":"W4389519305","doi":"10.18653/v1/2023.emnlp-main.160","title":"The Skipped Beat: A Study of Sociopragmatic Understanding in LLMs for 64 Languages","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Sparrow; Computer science; Benchmark (surveying); Meaning (existential); Spoken language; Scripting language; Task (project management); Natural language processing; Artificial intelligence; Psychology; Programming language; Geography","score_opus":0.08510672602283446,"score_gpt":0.3370588224076921,"score_spread":0.2519520963848576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519305","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8718883,0.009259558,0.053617153,0.0019150577,0.0007314334,0.0004509769,0.019411387,0.02520641,0.0175198],"genre_scores_gemma":[0.8842978,0.0008218745,0.040886506,0.0009618299,0.00015753774,0.00041290105,0.06276456,0.0015824855,0.0081144925],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9975228,0.0010723734,0.00017308697,0.00078684714,0.0002731728,0.00017171429],"domain_scores_gemma":[0.99233055,0.005357261,0.00024045711,0.0011165757,0.0006181124,0.00033697506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034933763,0.0024009186,0.001101236,0.0016389997,0.0009409142,0.0019078827,0.002861953,0.0020432002,0.0051662163],"category_scores_gemma":[0.016423717,0.0005949992,0.0014304455,0.0012186761,0.0010640546,0.0055730874,0.0029361031,0.0036873657,0.003418668],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031574785,0.0015120634,0.042153396,0.0032617114,0.0012139381,0.002059674,0.0042060725,0.1654101,0.016698234,0.0059049907,0.10835824,0.64606416],"study_design_scores_gemma":[0.00033144554,0.0014217816,0.02132162,0.00025235637,0.00026213023,0.001050632,0.0029121693,0.9148727,0.014526975,0.011981665,0.030877214,0.00018923129],"about_ca_topic_score_codex":0.013496703,"about_ca_topic_score_gemma":0.021329887,"teacher_disagreement_score":0.013496703,"about_ca_system_score_codex":0.0016257447,"about_ca_system_score_gemma":0.0012107528,"threshold_uncertainty_score":0.026836276},"labels":[],"label_agreement":null},{"id":"W4389519325","doi":"10.18653/v1/2023.findings-emnlp.84","title":"Can ChatGPT Assess Human Personalities? A General Evaluation Framework","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Research Foundation Singapore","keywords":"Personality psychology; Correctness; Computer science; Consistency (knowledge bases); Robustness (evolution); Psychology; Social psychology; Personality; Artificial intelligence; Algorithm","score_opus":0.14425137894513812,"score_gpt":0.3721872873750639,"score_spread":0.22793590842992578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519325","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24299137,0.0012254164,0.7278375,0.0017197041,0.0002478147,0.0032522324,0.0029457684,0.0064825113,0.013297768],"genre_scores_gemma":[0.8345131,0.00016342336,0.1587919,0.00037295473,0.00012119371,0.002747618,0.0017908213,0.00028470217,0.0012142353],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.93333733,0.052975465,0.0028476336,0.0040814225,0.0059433435,0.0008148238],"domain_scores_gemma":[0.83925617,0.12639993,0.00689361,0.014481944,0.010507743,0.0024605724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.047071308,0.0020061687,0.0011397348,0.003180997,0.00084558217,0.0034292422,0.0017563499,0.0021531077,0.0033809044],"category_scores_gemma":[0.18340123,0.00042764877,0.0011721943,0.0019729873,0.0019309472,0.0057371273,0.0035533193,0.0021297247,0.0009953713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006863373,0.0018756798,0.13515972,0.003118885,0.0019293299,0.00037818853,0.0069448394,0.09900729,0.024188187,0.037274025,0.015756834,0.6675037],"study_design_scores_gemma":[0.0006573058,0.0057496284,0.05744288,0.0005728521,0.0008932889,0.00055944116,0.0028757178,0.8040806,0.031875443,0.080210306,0.014675123,0.0004074013],"about_ca_topic_score_codex":0.002156244,"about_ca_topic_score_gemma":0.0016619571,"teacher_disagreement_score":0.047071308,"about_ca_system_score_codex":0.0015022977,"about_ca_system_score_gemma":0.0018433032,"threshold_uncertainty_score":0.24893987},"labels":[],"label_agreement":null},{"id":"W4389519370","doi":"10.18653/v1/2023.emnlp-main.920","title":"Unveiling the Essence of Poetry: Introducing a Comprehensive Dataset and Benchmark for Poem Summarization","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Stornoway Diamond (Canada); York University","funders":"","keywords":"Automatic summarization; Computer science; Meaning (existential); Natural language processing; Poetry; Task (project management); Literal and figurative language; Interpretation (philosophy); Artificial intelligence; Field (mathematics); Natural language; Natural language generation; Natural (archaeology); Benchmark (surveying); Linguistics; History; Epistemology; Programming language; Philosophy","score_opus":0.032744166450579815,"score_gpt":0.28308923003006053,"score_spread":0.2503450635794807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519370","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27848524,0.01325477,0.0802181,0.0053225774,0.0032604104,0.002712523,0.54280204,0.022136996,0.051807396],"genre_scores_gemma":[0.12094369,0.0013546268,0.08152149,0.0004951316,0.0004703227,0.0015156041,0.78491336,0.0005713878,0.008214348],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985896,0.00039996748,0.0002054071,0.00030689043,0.0003932965,0.00010481847],"domain_scores_gemma":[0.9977062,0.00070742547,0.00019832388,0.0005350212,0.00062749273,0.00022555918],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001300081,0.0010065909,0.0005804863,0.004401751,0.0014718713,0.0016516589,0.001704118,0.0015733375,0.004374878],"category_scores_gemma":[0.005943694,0.00017812806,0.000767608,0.004053928,0.00094395695,0.0023634161,0.002070282,0.0017331832,0.0052881343],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008058341,0.0011961919,0.013408304,0.0032032519,0.00014734242,0.0010270168,0.0022452204,0.006832715,0.0110077,0.010459818,0.6367343,0.31293228],"study_design_scores_gemma":[0.0003619162,0.00049043034,0.0432172,0.00048381704,0.00010531126,0.0018211197,0.004608314,0.043945316,0.017843315,0.0131442575,0.87381417,0.00016481426],"about_ca_topic_score_codex":0.0027374367,"about_ca_topic_score_gemma":0.0077420743,"teacher_disagreement_score":0.004401751,"about_ca_system_score_codex":0.00086962053,"about_ca_system_score_gemma":0.00083320093,"threshold_uncertainty_score":0.014635444},"labels":[],"label_agreement":null},{"id":"W4389519395","doi":"10.18653/v1/2023.emnlp-main.195","title":"Pushdown Layers: Encoding Recursive Structure in Transformer Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Center for Evolutionary and Theoretical Immunology; Canadian Institute for Advanced Research","keywords":"Computer science; Parsing; Transformer; Stack (abstract data type); Artificial intelligence; Algorithm; Theoretical computer science; Programming language","score_opus":0.025703243846209034,"score_gpt":0.2635919863223192,"score_spread":0.2378887424761102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519395","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038562812,0.00048376765,0.9383884,0.0003319225,0.0001229291,0.00011319965,0.0015364166,0.017122837,0.003337696],"genre_scores_gemma":[0.6863447,0.0005935683,0.29619083,0.00053738925,0.00008283792,0.00035870328,0.0043807575,0.0026672625,0.008843911],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996412,0.00009472302,0.000025511139,0.00013048349,0.000063051615,0.000045020322],"domain_scores_gemma":[0.998708,0.0007864541,0.00005757422,0.0002324028,0.00016274276,0.000052708772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00096857786,0.0012783337,0.00052958773,0.0007355437,0.00037704635,0.0013180847,0.0015635117,0.00085458823,0.0055866363],"category_scores_gemma":[0.0050714235,0.0006853823,0.0010652595,0.0006096157,0.0006550262,0.003497622,0.0017917781,0.0022081619,0.0026618266],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00057827355,0.00015218394,0.004299123,0.0003745272,0.00020902355,0.0004371874,0.0007434844,0.33422878,0.03008236,0.040280066,0.018265825,0.5703493],"study_design_scores_gemma":[0.000023410357,0.000039931885,0.00023781546,0.000018754989,0.000034707322,0.0000459285,0.000034304474,0.96228933,0.007392883,0.026959546,0.0029040554,0.000019331706],"about_ca_topic_score_codex":0.00630812,"about_ca_topic_score_gemma":0.016039025,"teacher_disagreement_score":0.00630812,"about_ca_system_score_codex":0.0009338008,"about_ca_system_score_gemma":0.0010316382,"threshold_uncertainty_score":0.018689096},"labels":[],"label_agreement":null},{"id":"W4389519409","doi":"10.18653/v1/2023.emnlp-main.16","title":"GPTAraEval: A Comprehensive Evaluation of ChatGPT on Arabic NLP","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Arabic; Modern Standard Arabic; Computer science; Natural language processing; Focus (optics); Transformative learning; Artificial intelligence; Software deployment; Linguistics; Psychology; Software engineering","score_opus":0.13639498830704402,"score_gpt":0.3458940014193512,"score_spread":0.2094990131123072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519409","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37800512,0.009069833,0.23945698,0.00467458,0.0024656418,0.0029635555,0.05967126,0.25399265,0.049700323],"genre_scores_gemma":[0.5803208,0.0019886317,0.22464794,0.0014705156,0.00028525892,0.0015486918,0.16809066,0.008238233,0.013409299],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9948778,0.0025642794,0.00036345673,0.0010686715,0.0009400328,0.00018568496],"domain_scores_gemma":[0.98589724,0.009670503,0.00024167648,0.0018607853,0.0017145986,0.00061524974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005579563,0.0024867048,0.0011851958,0.0022367672,0.0013499735,0.002550181,0.0033870088,0.0022764553,0.009981016],"category_scores_gemma":[0.026118685,0.000513944,0.0013574401,0.0017277317,0.00095796864,0.0048535652,0.0037138192,0.0032259815,0.005967899],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027484843,0.0014449212,0.015041438,0.0067171273,0.0014525685,0.0019082811,0.003869497,0.13874799,0.015972205,0.0067279553,0.26988542,0.5354841],"study_design_scores_gemma":[0.00060843164,0.001046873,0.009988479,0.0005108398,0.00034453926,0.0010019945,0.0021963625,0.8498735,0.014361883,0.008014932,0.11183257,0.00021963372],"about_ca_topic_score_codex":0.02208818,"about_ca_topic_score_gemma":0.024740774,"teacher_disagreement_score":0.02208818,"about_ca_system_score_codex":0.0018127545,"about_ca_system_score_gemma":0.0022901772,"threshold_uncertainty_score":0.043919206},"labels":[],"label_agreement":null},{"id":"W4389519413","doi":"10.18653/v1/2023.findings-emnlp.86","title":"Large Language Models Know Your Contextual Search Intent: A Prompting Framework for Conversational Search","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Renmin University of China; National Natural Science Foundation of China","keywords":"Computer science; Leverage (statistics); Conversation; Language model; Robustness (evolution); Human–computer interaction; Natural language processing; Artificial intelligence; Information retrieval; World Wide Web; Data science; Psychology; Communication","score_opus":0.10600733513719265,"score_gpt":0.3566575900217333,"score_spread":0.2506502548845406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519413","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023259424,0.0007436211,0.9663902,0.00090378925,0.000060556857,0.00016636399,0.0007837783,0.0058404664,0.0018517759],"genre_scores_gemma":[0.5255285,0.00046075144,0.4676085,0.00037504436,0.00014701538,0.00039035987,0.0022642056,0.00045401888,0.0027716917],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99882835,0.0006361123,0.000072686744,0.00024025212,0.00015880879,0.00006381189],"domain_scores_gemma":[0.9970716,0.0018235366,0.00017404713,0.0004197335,0.00036550188,0.00014554974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019776658,0.000988085,0.000634438,0.0012421617,0.00061606517,0.0013039383,0.0013178989,0.0009926249,0.0024684467],"category_scores_gemma":[0.0086591095,0.00040526208,0.0010931159,0.0007729001,0.00064017606,0.0028833996,0.0014804518,0.0018843802,0.0014356128],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090072816,0.00043851204,0.007288398,0.0009245045,0.00021034142,0.0005718223,0.004380095,0.2683261,0.036656916,0.061064497,0.02291136,0.5963267],"study_design_scores_gemma":[0.00003108206,0.00008133258,0.00042300334,0.000025289306,0.00004045959,0.00009504082,0.00020889919,0.9640265,0.0030894564,0.02617709,0.0057693636,0.000032530315],"about_ca_topic_score_codex":0.0067615933,"about_ca_topic_score_gemma":0.010338143,"teacher_disagreement_score":0.0067615933,"about_ca_system_score_codex":0.000916071,"about_ca_system_score_gemma":0.0017251077,"threshold_uncertainty_score":0.013444424},"labels":[],"label_agreement":null},{"id":"W4389519419","doi":"10.18653/v1/2023.emnlp-main.37","title":"Knowledge Graph Compression Enhances Diverse Commonsense Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Commonsense knowledge; Computer science; Commonsense reasoning; Knowledge graph; Artificial intelligence; Graph; Context (archaeology); Task (project management); Natural language processing; Theoretical computer science; Machine learning; Knowledge-based systems","score_opus":0.09308819212558869,"score_gpt":0.3087466646658838,"score_spread":0.2156584725402951,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519419","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10797435,0.0004236429,0.8803434,0.00090681994,0.0000514693,0.0001886074,0.0005345918,0.0025389718,0.0070381863],"genre_scores_gemma":[0.73762345,0.00031208716,0.2568256,0.0003023612,0.000060362538,0.00019127494,0.0017048905,0.0003464987,0.0026334526],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993641,0.00019589522,0.000025988396,0.00018217959,0.00018018614,0.00005181137],"domain_scores_gemma":[0.9959111,0.003028187,0.00014739315,0.0005276086,0.00028196504,0.000103830695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010220488,0.0007509835,0.0005973655,0.0012552554,0.0004410149,0.0009465173,0.0013539535,0.0011344521,0.0030763522],"category_scores_gemma":[0.007940157,0.0002968521,0.00088557746,0.000985639,0.0009047877,0.0027580524,0.0017434392,0.0013007326,0.0005034769],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023048728,0.00030275388,0.0026460486,0.00024496665,0.0000729005,0.00056450936,0.00056908827,0.6609557,0.012253321,0.059285272,0.0048163566,0.25805855],"study_design_scores_gemma":[0.000017666925,0.000021204758,0.00020458149,0.00001270202,0.000016689646,0.00006546973,0.000037603237,0.9608444,0.0033409516,0.033920728,0.0015110688,0.0000068932427],"about_ca_topic_score_codex":0.0028582432,"about_ca_topic_score_gemma":0.0047248853,"teacher_disagreement_score":0.0030763522,"about_ca_system_score_codex":0.0008788154,"about_ca_system_score_gemma":0.0009294978,"threshold_uncertainty_score":0.010291398},"labels":[],"label_agreement":null},{"id":"W4389519420","doi":"10.18653/v1/2023.emnlp-main.142","title":"A Diachronic Analysis of Paradigm Shifts in NLP Research: When, How, and Why?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Deutsche Forschungsgemeinschaft","keywords":"Leverage (statistics); Computer science; Inference; Causal inference; Artificial intelligence; Field (mathematics); Data science; Natural language processing; Machine learning; Mathematics","score_opus":0.1503477852821135,"score_gpt":0.3498817030182792,"score_spread":0.1995339177361657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519420","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23183303,0.007778958,0.7144098,0.017203992,0.00032115448,0.00051594956,0.002788338,0.00082095905,0.024327775],"genre_scores_gemma":[0.80937284,0.0022609348,0.18305436,0.00081393303,0.0002606823,0.00045481234,0.0015560138,0.00020927285,0.002017094],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9888782,0.0053855325,0.0008573183,0.002328077,0.0021969944,0.0003538297],"domain_scores_gemma":[0.8674789,0.10350247,0.013163877,0.0074802428,0.0070873667,0.0012870315],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.021290362,0.0005790303,0.00052092824,0.015283661,0.0035512098,0.0073931147,0.0015412575,0.0014043428,0.0033616326],"category_scores_gemma":[0.10137575,0.0005817076,0.0008684935,0.014766478,0.0057782573,0.018844454,0.0039333403,0.0038278822,0.0006455634],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025087956,0.00018650464,0.13255246,0.0013032756,0.00019931464,0.0005406053,0.022177352,0.007257249,0.007484885,0.51961476,0.004889836,0.30354282],"study_design_scores_gemma":[0.000038070917,0.00008750488,0.06154862,0.00065853616,0.0001247407,0.00075612724,0.012034459,0.06112025,0.0049951063,0.810462,0.04803076,0.00014384692],"about_ca_topic_score_codex":0.0031978714,"about_ca_topic_score_gemma":0.0048022172,"teacher_disagreement_score":0.98471636,"about_ca_system_score_codex":0.0042118067,"about_ca_system_score_gemma":0.0044497545,"threshold_uncertainty_score":0.11259556},"labels":[],"label_agreement":null},{"id":"W4389519429","doi":"10.18653/v1/2023.emnlp-main.715","title":"mAggretriever: A Simple yet Effective Approach to Zero-Shot Multilingual Dense Retrieval","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canada First Research Excellence Fund","keywords":"Computer science; Inference; Transformer; Natural language processing; Artificial intelligence; Task (project management); Language model; Simple (philosophy); Training set; Information retrieval","score_opus":0.0394008374428256,"score_gpt":0.29520197189589986,"score_spread":0.2558011344530743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519429","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013295748,0.0009901493,0.94413483,0.00017104385,0.00017648752,0.00025835945,0.00073591404,0.03680479,0.0034326536],"genre_scores_gemma":[0.21480714,0.0006161238,0.76348436,0.0007655173,0.00020060914,0.00044902976,0.005306071,0.0030068909,0.011364225],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998966,0.00018457106,0.000064476546,0.00038421186,0.00025273315,0.0001481624],"domain_scores_gemma":[0.99875295,0.00039318006,0.000056361492,0.00048067214,0.0002388747,0.000078024466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016631007,0.0016101205,0.0015779381,0.0016523249,0.0007263339,0.0014907725,0.004679477,0.0013948459,0.010785295],"category_scores_gemma":[0.004835443,0.00093958346,0.0011973543,0.0013073108,0.0008548615,0.004300531,0.0042413133,0.0017739023,0.007937134],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047811284,0.00035868346,0.0010673772,0.000480213,0.0002389651,0.0002672008,0.00029083647,0.05426925,0.0436623,0.010089407,0.03652818,0.8522695],"study_design_scores_gemma":[0.00009451126,0.00021242333,0.00046990177,0.000029602448,0.00007508732,0.00037125807,0.00012557073,0.94048595,0.025363976,0.015713964,0.016988406,0.00006931375],"about_ca_topic_score_codex":0.010572714,"about_ca_topic_score_gemma":0.023018898,"teacher_disagreement_score":0.010785295,"about_ca_system_score_codex":0.0008957027,"about_ca_system_score_gemma":0.0019025711,"threshold_uncertainty_score":0.03608036},"labels":[],"label_agreement":null},{"id":"W4389519446","doi":"10.18653/v1/2023.emnlp-main.782","title":"The CoT Collection: Improving Zero-shot and Few-shot Learning of Language Models via Chain-of-Thought Fine-Tuning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Benchmark (surveying); Task (project management); Artificial intelligence","score_opus":0.03176144178090691,"score_gpt":0.2597585747241698,"score_spread":0.2279971329432629,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519446","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26431474,0.011102606,0.6139295,0.0017230001,0.001251169,0.00088452006,0.007359681,0.08771502,0.011719792],"genre_scores_gemma":[0.62461054,0.000770688,0.33601692,0.0021119742,0.00027471015,0.0006657112,0.02551076,0.0018003968,0.008238343],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9975858,0.0007872398,0.00015539215,0.0009280775,0.00035474123,0.00018868539],"domain_scores_gemma":[0.9948186,0.003274382,0.00017903163,0.00096920424,0.00049825385,0.0002604696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029586018,0.0028482461,0.0016201902,0.0012649015,0.0007439297,0.0018563308,0.0044396226,0.0027871656,0.0048680967],"category_scores_gemma":[0.013681112,0.0008604349,0.002165729,0.0009388112,0.0010515847,0.00428789,0.0024974535,0.0051741186,0.0030024026],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012388977,0.0012536885,0.0077316477,0.0013878111,0.00063531206,0.0005586234,0.0004570266,0.14997044,0.023232883,0.004148452,0.047838308,0.7615469],"study_design_scores_gemma":[0.00028540078,0.0005195089,0.001462697,0.00007806845,0.0001297064,0.00019130776,0.00015250173,0.9698706,0.009998414,0.010118867,0.00711545,0.000077385455],"about_ca_topic_score_codex":0.009164261,"about_ca_topic_score_gemma":0.020207806,"teacher_disagreement_score":0.009164261,"about_ca_system_score_codex":0.0014041048,"about_ca_system_score_gemma":0.002366035,"threshold_uncertainty_score":0.018221855},"labels":[],"label_agreement":null},{"id":"W4389519516","doi":"10.18653/v1/2023.emnlp-main.955","title":"Anaphor Assisted Document-Level Relation Extraction","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment","keywords":"Computer science; Relationship extraction; Focus (optics); Graph; Natural language processing; Relation (database); Sentence; Artificial intelligence; Information retrieval; Information extraction; Theoretical computer science; Data mining","score_opus":0.08087274730028714,"score_gpt":0.31445749031028564,"score_spread":0.23358474300999849,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519516","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021878527,0.0064567532,0.92861795,0.00078098074,0.00030433948,0.0005851133,0.010562312,0.024374204,0.0064397776],"genre_scores_gemma":[0.14323337,0.00270304,0.79852545,0.00053091056,0.00035834208,0.00044156978,0.04157556,0.000680786,0.011950953],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977633,0.000363131,0.0002456943,0.0007877972,0.0006580208,0.00018204586],"domain_scores_gemma":[0.9972652,0.0008778749,0.00027731914,0.0007703177,0.0007201066,0.00008921013],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017880575,0.0020135888,0.0018068134,0.007660776,0.0013119078,0.00199493,0.0028173432,0.002082712,0.0036934724],"category_scores_gemma":[0.0050302274,0.00058191834,0.0020829083,0.007215637,0.00053453003,0.005196502,0.0021315424,0.0020678688,0.0053415303],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028229062,0.00038493954,0.0033740192,0.0010194257,0.0002437541,0.0008781494,0.00047650788,0.01087569,0.031156683,0.01477199,0.084968075,0.8515685],"study_design_scores_gemma":[0.00013454731,0.00020365953,0.006931143,0.00017845476,0.0005556933,0.0034467576,0.00056545006,0.7072521,0.078183666,0.04672992,0.15560335,0.00021536306],"about_ca_topic_score_codex":0.005495077,"about_ca_topic_score_gemma":0.011863055,"teacher_disagreement_score":0.007660776,"about_ca_system_score_codex":0.0008520636,"about_ca_system_score_gemma":0.002565939,"threshold_uncertainty_score":0.012355864},"labels":[],"label_agreement":null},{"id":"W4389519518","doi":"10.18653/v1/2023.emnlp-main.844","title":"Aligning Large Language Models through Synthetic Feedback","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Reinforcement learning; Language model; Artificial intelligence; Quality (philosophy); Machine learning","score_opus":0.03652306064352353,"score_gpt":0.2776569441957399,"score_spread":0.24113388355221635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519518","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20281222,0.00096973457,0.76813537,0.0013470873,0.00029197097,0.00029359324,0.0015258234,0.01884747,0.005776727],"genre_scores_gemma":[0.84632844,0.0001948521,0.14482406,0.0004534242,0.00007303345,0.00046533547,0.0030465696,0.0009452489,0.0036689257],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9972422,0.0016825042,0.00008156432,0.00059713755,0.00026937656,0.00012719377],"domain_scores_gemma":[0.99169606,0.0062133567,0.00035095273,0.00085929653,0.00064409146,0.00023626575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003399259,0.001555664,0.0007506529,0.0005290116,0.00042441665,0.001070438,0.0018660297,0.0016507264,0.0038922264],"category_scores_gemma":[0.02134588,0.00056247425,0.0006732002,0.0005247626,0.0009783715,0.0022176057,0.0018442974,0.0023876226,0.0016671591],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005890268,0.00029348902,0.0036556139,0.00042596873,0.000099716955,0.00035097057,0.000536404,0.8500321,0.011963679,0.0081308,0.009584054,0.114338204],"study_design_scores_gemma":[0.00003403852,0.0000740865,0.0001974201,0.000013489374,0.000008813563,0.000031560234,0.000041857154,0.99144286,0.003123114,0.0036443272,0.0013753022,0.000013123807],"about_ca_topic_score_codex":0.004845713,"about_ca_topic_score_gemma":0.007662251,"teacher_disagreement_score":0.004845713,"about_ca_system_score_codex":0.0010586063,"about_ca_system_score_gemma":0.0013489954,"threshold_uncertainty_score":0.017977178},"labels":[],"label_agreement":null},{"id":"W4389519573","doi":"10.18653/v1/2023.findings-emnlp.98","title":"Dolphin: A Challenging and Diverse Benchmark for Arabic NLG","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Computer science; Benchmark (surveying); Automatic summarization; Natural language processing; Modular design; Artificial intelligence; Set (abstract data type); Modern Standard Arabic; Generalization; Arabic; Linguistics; Programming language","score_opus":0.0487968856808134,"score_gpt":0.2677313356933869,"score_spread":0.21893445001257353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519573","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33376002,0.018521428,0.20773731,0.010590664,0.004414169,0.0043419763,0.20185801,0.08501227,0.13376418],"genre_scores_gemma":[0.32501447,0.0024722978,0.21213078,0.0020254704,0.0005265559,0.0028833554,0.43185383,0.0062536937,0.01683944],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9923463,0.0036753577,0.0007330459,0.0013921909,0.0014469366,0.00040620955],"domain_scores_gemma":[0.9871157,0.006424061,0.00040009458,0.0028265843,0.0025739914,0.0006596261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008658279,0.0029838984,0.001277803,0.0055981586,0.0028521167,0.0037905697,0.0036262465,0.003100671,0.009809377],"category_scores_gemma":[0.030062655,0.00051242026,0.0014702904,0.0054299505,0.0017202769,0.005207325,0.0050315065,0.002600463,0.007268904],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015483105,0.0012323248,0.010694628,0.0038147017,0.0005234099,0.00096136425,0.0022159547,0.11060836,0.008215095,0.020675782,0.47328225,0.3662279],"study_design_scores_gemma":[0.0008992854,0.000766024,0.011814076,0.00081745966,0.00018944751,0.0015141953,0.0031750537,0.502055,0.027055165,0.04599812,0.40546054,0.00025564551],"about_ca_topic_score_codex":0.017805157,"about_ca_topic_score_gemma":0.020122996,"teacher_disagreement_score":0.017805157,"about_ca_system_score_codex":0.002942663,"about_ca_system_score_gemma":0.0029742236,"threshold_uncertainty_score":0.045789897},"labels":[],"label_agreement":null},{"id":"W4389519576","doi":"10.18653/v1/2023.emnlp-industry.33","title":"Building Real-World Meeting Summarization Systems using Large Language Models: A Practical Perspective","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Stornoway Diamond (Canada)","funders":"","keywords":"Automatic summarization; Computer science; Open source; Perspective (graphical); Artificial intelligence; Programming language; Software","score_opus":0.06953591968372314,"score_gpt":0.35937517655501566,"score_spread":0.2898392568712925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519576","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032072235,0.001277604,0.8960567,0.002495893,0.00029801656,0.00052215124,0.0028361988,0.059414554,0.0050265845],"genre_scores_gemma":[0.2965634,0.0007596294,0.6823076,0.0008298631,0.00019651663,0.00049238093,0.012571555,0.0024986267,0.0037804616],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99545825,0.0025133227,0.0003131035,0.0007981741,0.0007228294,0.00019440075],"domain_scores_gemma":[0.98924595,0.0059479238,0.0005295442,0.0024835565,0.0014207822,0.00037223785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052350536,0.00167857,0.0011672905,0.0013906354,0.00080411654,0.0036737218,0.0029081244,0.0015927365,0.005491292],"category_scores_gemma":[0.02761169,0.0008159244,0.0013784432,0.0014455203,0.00074583886,0.007716022,0.0029117272,0.0022655872,0.006146819],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015273204,0.00085992436,0.007907891,0.0023245465,0.000673445,0.000726952,0.0028132794,0.24832329,0.045398846,0.029758481,0.08432468,0.5753613],"study_design_scores_gemma":[0.0000980711,0.00026738233,0.00086844055,0.00008789955,0.000111090965,0.00020889555,0.0009892326,0.9190692,0.014640456,0.023990318,0.0395721,0.00009692844],"about_ca_topic_score_codex":0.005329194,"about_ca_topic_score_gemma":0.008512265,"teacher_disagreement_score":0.005491292,"about_ca_system_score_codex":0.0013075056,"about_ca_system_score_gemma":0.0017094144,"threshold_uncertainty_score":0.027686},"labels":[],"label_agreement":null},{"id":"W4389519586","doi":"10.18653/v1/2023.emnlp-main.971","title":"MQuAKE: Assessing Knowledge Editing in Language Models via Multi-Hop Questions","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research; National Science Foundation","keywords":"Computer science; Recall; Language model; Question answering; Margin (machine learning); Natural language processing; Hop (telecommunications); Artificial intelligence; Machine learning; Cognitive psychology; Psychology","score_opus":0.06529112178130436,"score_gpt":0.3471593056174225,"score_spread":0.2818681838361181,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519586","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61410385,0.005460775,0.32398933,0.0015484623,0.0005868205,0.001661001,0.010284572,0.031373654,0.010991553],"genre_scores_gemma":[0.804905,0.00047594236,0.17624702,0.00059268257,0.000104716804,0.00074350816,0.013117599,0.0007929933,0.0030205846],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98849934,0.006745114,0.00082582154,0.0019463422,0.0017447232,0.00023862973],"domain_scores_gemma":[0.8620288,0.12276412,0.0034404069,0.0073099993,0.0031551984,0.0013015101],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013767496,0.0021038523,0.0010475701,0.0028246588,0.00074331247,0.0030214235,0.0028385483,0.0036101285,0.004119713],"category_scores_gemma":[0.118785806,0.0005633257,0.0012667098,0.0011978392,0.0008768636,0.0077374345,0.004348249,0.0036228043,0.0016212781],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004110587,0.0031443546,0.07030745,0.0055759847,0.0026409915,0.0007052181,0.0078016473,0.22265205,0.026736833,0.0091499835,0.033994943,0.61318004],"study_design_scores_gemma":[0.00041849108,0.002215125,0.02784866,0.00033127744,0.0005424622,0.0006770639,0.0017078294,0.8882799,0.026990687,0.028997371,0.021674177,0.00031704095],"about_ca_topic_score_codex":0.0044052484,"about_ca_topic_score_gemma":0.006128645,"teacher_disagreement_score":0.013767496,"about_ca_system_score_codex":0.001006835,"about_ca_system_score_gemma":0.0011706494,"threshold_uncertainty_score":0.07281035},"labels":[],"label_agreement":null},{"id":"W4389519588","doi":"10.18653/v1/2023.emnlp-main.134","title":"MAGNIFICo: Evaluating the In-Context Learning Ability of Large Language Models to Generalize to Novel Interpretations","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Context (archaeology); Computer science; Parsing; Natural language processing; Artificial intelligence; Suite; Natural language; Cognitive science; Linguistics; Cognitive psychology; Psychology; History","score_opus":0.08345484632105207,"score_gpt":0.36307393927965914,"score_spread":0.2796190929586071,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519588","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77183104,0.0037420709,0.16803934,0.0020457257,0.0005882182,0.001254793,0.0061779427,0.029257935,0.017062947],"genre_scores_gemma":[0.8203873,0.0005862106,0.1624375,0.000601658,0.00016934128,0.000728401,0.0106426515,0.0013242727,0.0031225784],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99617815,0.0024256925,0.00019701633,0.00067250064,0.00036277543,0.00016389962],"domain_scores_gemma":[0.9723288,0.023213489,0.00064888614,0.00236286,0.0009110801,0.0005349113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077070277,0.002720045,0.00081452925,0.0014830035,0.0007209583,0.0017065338,0.0023770605,0.0024668844,0.003418676],"category_scores_gemma":[0.031046899,0.00057585357,0.0011897292,0.0006373351,0.0010834298,0.004141639,0.003054385,0.0029595783,0.0010501686],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038706416,0.0037270787,0.033464897,0.0030894238,0.0013508437,0.0007188453,0.0031673387,0.4150163,0.018071493,0.013213072,0.047031656,0.45727843],"study_design_scores_gemma":[0.0003282537,0.0011069549,0.0050012255,0.00011258751,0.00013852504,0.00022750224,0.0005359857,0.96298105,0.010015242,0.011716005,0.007747379,0.00008918067],"about_ca_topic_score_codex":0.006670925,"about_ca_topic_score_gemma":0.009575141,"teacher_disagreement_score":0.0077070277,"about_ca_system_score_codex":0.0011276905,"about_ca_system_score_gemma":0.0013270336,"threshold_uncertainty_score":0.040759146},"labels":[],"label_agreement":null},{"id":"W4389519609","doi":"10.18653/v1/2023.emnlp-main.144","title":"Syntactic Substitutability as Unsupervised Dependency Syntax","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Natural Sciences and Engineering Research Council of Canada; Universities Space Research Association","keywords":"Parsing; Computer science; Natural language processing; Syntax; Dependency (UML); Artificial intelligence; Dependency grammar; Property (philosophy); Precision and recall; Natural language","score_opus":0.03197503169749673,"score_gpt":0.2716625714997887,"score_spread":0.23968753980229196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519609","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03569816,0.00024358091,0.9553306,0.00030711974,0.00004529433,0.00008836238,0.0011208628,0.004521809,0.0026442576],"genre_scores_gemma":[0.628554,0.00037053163,0.3573112,0.00029297546,0.00010354457,0.00028482702,0.005298181,0.0014602897,0.006324512],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9987062,0.00037239297,0.000074846575,0.0005650543,0.00020278752,0.00007874373],"domain_scores_gemma":[0.9957504,0.002444365,0.0003892041,0.0008994794,0.00043890375,0.00007750238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017921976,0.0009379252,0.0005556835,0.0025102932,0.00065549545,0.0013424251,0.0015702688,0.0009301711,0.0036477612],"category_scores_gemma":[0.0056813797,0.00071984436,0.0015088501,0.001962725,0.000958106,0.0035276448,0.0016151582,0.0019919768,0.0015463539],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002944614,0.00020016127,0.015393651,0.00056561205,0.0004045591,0.00053559756,0.0017198687,0.0827062,0.04342659,0.17182325,0.026869388,0.6560606],"study_design_scores_gemma":[0.000021709007,0.000044656925,0.0048094303,0.000052328516,0.00013076473,0.00034919346,0.0001568294,0.8019848,0.01494123,0.1638317,0.013619841,0.000057561974],"about_ca_topic_score_codex":0.0036325837,"about_ca_topic_score_gemma":0.007924183,"teacher_disagreement_score":0.0036477612,"about_ca_system_score_codex":0.0011758158,"about_ca_system_score_gemma":0.0017142026,"threshold_uncertainty_score":0.012203038},"labels":[],"label_agreement":null},{"id":"W4389519614","doi":"10.18653/v1/2023.emnlp-main.1040","title":"JASMINE: Arabic GPT Models for Few-Shot Learning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Arabic; Islam; Computer science; Natural language processing; Artificial intelligence; Linguistics; History; Philosophy; Archaeology","score_opus":0.08778298209702809,"score_gpt":0.29335306340552875,"score_spread":0.20557008130850066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519614","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011773221,0.0011900078,0.92880756,0.0005798824,0.0003309673,0.00017559278,0.0041607665,0.04731236,0.005669635],"genre_scores_gemma":[0.21635172,0.0009542903,0.73636615,0.0007958929,0.00020708617,0.0007099117,0.021252837,0.0047804792,0.018581651],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99960774,0.00009692087,0.00002870831,0.00013790072,0.000093311406,0.00003542306],"domain_scores_gemma":[0.9990133,0.00047239164,0.000035555408,0.00019097135,0.00022241063,0.000065202854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011459461,0.0012198031,0.0010633097,0.0010654913,0.0007141747,0.0019378158,0.0033846002,0.0020733676,0.017212681],"category_scores_gemma":[0.005954348,0.0006352887,0.0013999895,0.0008826291,0.00035110756,0.003113645,0.0020564653,0.0030021842,0.010096494],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008982438,0.0003315591,0.0016184486,0.00039107486,0.00029945516,0.0003098343,0.00023701214,0.25577608,0.00525484,0.019095743,0.08142474,0.63436306],"study_design_scores_gemma":[0.00003646783,0.000029716668,0.00011852335,0.000018162507,0.00001910702,0.000044491127,0.000022890757,0.97778666,0.0016140257,0.013274487,0.0070209643,0.000014453552],"about_ca_topic_score_codex":0.010988968,"about_ca_topic_score_gemma":0.01802793,"teacher_disagreement_score":0.017212681,"about_ca_system_score_codex":0.0010360976,"about_ca_system_score_gemma":0.0010833849,"threshold_uncertainty_score":0.0575822},"labels":[],"label_agreement":null},{"id":"W4389519618","doi":"10.18653/v1/2023.emnlp-main.125","title":"Self-Influence Guided Data Reweighting for Language Model Pre-training","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Novelty; Artificial intelligence; Language model; Context (archaeology); Task (project management); Machine learning; Relevance (law); Sample (material); Training set; Stability (learning theory); Point (geometry); Data modeling; Natural language processing","score_opus":0.12640408359377966,"score_gpt":0.3563720179076763,"score_spread":0.22996793431389664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519618","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035892867,0.0007931038,0.9571903,0.00035003276,0.00017942848,0.00016020717,0.00018924245,0.0040720925,0.0011728107],"genre_scores_gemma":[0.47891095,0.00042145085,0.50997144,0.0006677943,0.0003767147,0.00067429716,0.0020820019,0.0020906013,0.0048047868],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99685836,0.0012843027,0.00023294693,0.0007846638,0.00062174594,0.00021795987],"domain_scores_gemma":[0.98864937,0.0066877366,0.0005424467,0.0021515056,0.0015665379,0.0004023206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006127829,0.0019492806,0.0018766337,0.0013528552,0.00095318747,0.0020745252,0.0029723237,0.0021903391,0.0029011448],"category_scores_gemma":[0.030583555,0.0009701553,0.0014978523,0.0009914034,0.0015960278,0.0038214058,0.0037798963,0.0053386963,0.002148709],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001071188,0.000711013,0.008740951,0.0006281666,0.0005040221,0.00028719698,0.0010667309,0.29467252,0.037234016,0.013786458,0.01171833,0.6295794],"study_design_scores_gemma":[0.000058131496,0.00018619295,0.0007459244,0.000042048418,0.000051906874,0.0000796713,0.00007641927,0.9735874,0.013201321,0.009242835,0.0027000809,0.000028097958],"about_ca_topic_score_codex":0.0021498266,"about_ca_topic_score_gemma":0.006009087,"teacher_disagreement_score":0.006127829,"about_ca_system_score_codex":0.00080891215,"about_ca_system_score_gemma":0.0016601644,"threshold_uncertainty_score":0.032407463},"labels":[],"label_agreement":null},{"id":"W4389519826","doi":"10.18653/v1/2023.findings-emnlp.245","title":"FinePrompt: Unveiling the Role of Finetuned Inductive Bias on Compositional Reasoning in GPT-4","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea","keywords":"Computer science; Leverage (statistics); Task (project management); Inductive bias; Inductive reasoning; Question answering; Artificial intelligence; Human–computer interaction; Natural language processing; Multi-task learning","score_opus":0.04108843627987746,"score_gpt":0.26247931949820996,"score_spread":0.2213908832183325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519826","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028957723,0.00021109857,0.9527691,0.0005941099,0.00012656224,0.00024475725,0.00040686087,0.014078032,0.0026117747],"genre_scores_gemma":[0.32668367,0.00018006122,0.6660444,0.0005763509,0.00010942471,0.0003448489,0.0013223765,0.0012864328,0.0034524247],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99714786,0.0013357716,0.0001583228,0.0006989015,0.0004954751,0.00016367831],"domain_scores_gemma":[0.9872362,0.008120116,0.00042127754,0.0028404782,0.0009877111,0.00039411633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004612308,0.0012256948,0.0008565848,0.00065613235,0.0007367497,0.0024863568,0.0023847588,0.0016861759,0.0070661334],"category_scores_gemma":[0.026145091,0.00058355427,0.0011536507,0.0005829036,0.001875257,0.006974253,0.004686947,0.003690864,0.0026726692],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016248749,0.0006382288,0.0062419237,0.0022014736,0.00015265626,0.00094638346,0.007501113,0.09109602,0.0835075,0.16883646,0.02129038,0.61596304],"study_design_scores_gemma":[0.00023172772,0.00034791825,0.00066231587,0.0001127146,0.000092901755,0.00039413138,0.00063389033,0.6905217,0.036539953,0.2437732,0.026614364,0.000075076096],"about_ca_topic_score_codex":0.0014534399,"about_ca_topic_score_gemma":0.0026695724,"teacher_disagreement_score":0.0070661334,"about_ca_system_score_codex":0.0009830152,"about_ca_system_score_gemma":0.00201083,"threshold_uncertainty_score":0.024392545},"labels":[],"label_agreement":null},{"id":"W4389519833","doi":"10.18653/v1/2023.findings-emnlp.219","title":"SWEET - Weakly Supervised Person Name Extraction for Fighting Human Trafficking","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Samsung; Canadian Institute for Advanced Research","keywords":"Benchmark (surveying); Pipeline (software); Computer science; Generalizability theory; Matching (statistics); Domain (mathematical analysis); Task (project management); Artificial intelligence; Labeled data; Sequence labeling; Machine learning; Data mining; Natural language processing; Mathematics; Engineering","score_opus":0.07735140164470757,"score_gpt":0.3088931297072098,"score_spread":0.23154172806250223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057827696,0.0010425923,0.89847016,0.00066692,0.00026190103,0.00028025248,0.0035656132,0.029691525,0.008193216],"genre_scores_gemma":[0.31844714,0.0005576527,0.6186494,0.0010517775,0.0003332759,0.000402373,0.031108754,0.0017523933,0.027697226],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984889,0.00039671388,0.000082713705,0.00058429013,0.00030782964,0.0001395712],"domain_scores_gemma":[0.99818015,0.00061429996,0.00018721087,0.0005694728,0.00036355288,0.00008538407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013438707,0.0016142665,0.0010726668,0.001959428,0.00096768694,0.0011584635,0.0015496608,0.0014840156,0.003707693],"category_scores_gemma":[0.003806951,0.0004603739,0.0010718069,0.001106983,0.00088221923,0.00304977,0.0019822924,0.0016691055,0.007528955],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006690409,0.0005578077,0.010828049,0.0005298842,0.0002247337,0.00075446226,0.00071779254,0.025952227,0.049830686,0.009348482,0.081984445,0.8186024],"study_design_scores_gemma":[0.00006645602,0.00022658132,0.0053029973,0.0001037074,0.00012673458,0.0011430657,0.00060879986,0.84970766,0.053453006,0.022110937,0.06706232,0.00008782064],"about_ca_topic_score_codex":0.0035852734,"about_ca_topic_score_gemma":0.0092316205,"teacher_disagreement_score":0.003707693,"about_ca_system_score_codex":0.00042461796,"about_ca_system_score_gemma":0.0014540241,"threshold_uncertainty_score":0.012403488},"labels":[],"label_agreement":null},{"id":"W4389519946","doi":"10.18653/v1/2023.findings-emnlp.300","title":"Verb Conjugation in Transformers Is Determined by Linear Encodings of Subject Number","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; National Science Foundation","keywords":"Verb; Subject (documents); Computer science; Transformer; Encoding (memory); Mathematics Subject Classification; Position (finance); Artificial intelligence; Layer (electronics); Linguistics; Natural language processing; Arithmetic; Speech recognition; Mathematics; Discrete mathematics; Philosophy; Physics","score_opus":0.01973994559083426,"score_gpt":0.269174178691716,"score_spread":0.24943423310088175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519946","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25483695,0.00012248405,0.7012023,0.00084818725,0.00014104525,0.00007574331,0.00045939913,0.004455682,0.03785828],"genre_scores_gemma":[0.95180506,0.00011428109,0.042028107,0.00014689214,0.000020869253,0.000037996328,0.00037104273,0.0005928932,0.004882922],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993268,0.00019570977,0.00004451115,0.00023388104,0.000111195084,0.00008794166],"domain_scores_gemma":[0.9978181,0.0010609411,0.00024205908,0.0005773757,0.00017304346,0.00012840368],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009501242,0.0005701869,0.0003868247,0.0004584265,0.00037914576,0.0031381394,0.00077722524,0.0005550539,0.008371241],"category_scores_gemma":[0.0066338535,0.0006040847,0.0006467816,0.00041587083,0.0017591144,0.0072768666,0.0020322322,0.0017277991,0.0020165145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078281556,0.00012779735,0.010142041,0.00037907306,0.00009037867,0.000437988,0.0031716945,0.01842511,0.14473933,0.5374826,0.0052355775,0.27898553],"study_design_scores_gemma":[0.00007119741,0.00026103863,0.00595467,0.00007770864,0.00014650475,0.00040905026,0.0007530929,0.21199909,0.07697494,0.68757087,0.015703931,0.000077876015],"about_ca_topic_score_codex":0.0013664211,"about_ca_topic_score_gemma":0.0015135976,"teacher_disagreement_score":0.008371241,"about_ca_system_score_codex":0.0008935937,"about_ca_system_score_gemma":0.0007684987,"threshold_uncertainty_score":0.028004587},"labels":[],"label_agreement":null},{"id":"W4389519982","doi":"10.18653/v1/2023.findings-emnlp.97","title":"On the Risk of Misinformation Pollution with Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Ministry of Education - Singapore; Bộ Giáo dục và Ðào tạo; Ministry of Education, India; National Science Foundation","keywords":"Misinformation; Obstacle; Harm; Computer science; Internet privacy; Risk analysis (engineering); Computer security; Data science; Psychology; Political science; Business; Social psychology","score_opus":0.015910636126625174,"score_gpt":0.2270961714336405,"score_spread":0.21118553530701534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519982","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38993034,0.0020880403,0.5903274,0.006153452,0.00022776308,0.00029658462,0.00041453174,0.002787529,0.0077742683],"genre_scores_gemma":[0.93520737,0.0003424722,0.062153615,0.0005770902,0.000116323616,0.00009335034,0.0002025248,0.00022122836,0.0010860394],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98341537,0.011737773,0.0004781472,0.0014276473,0.0024158903,0.0005250623],"domain_scores_gemma":[0.71443754,0.24752891,0.010537838,0.020133054,0.006085249,0.0012773726],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01830155,0.0012549158,0.0011844503,0.0018861513,0.0012611678,0.0035893451,0.0015260514,0.0025653655,0.001560135],"category_scores_gemma":[0.14235865,0.0008005427,0.0009146158,0.0011842726,0.0028398605,0.007800676,0.004034901,0.0038996506,0.0006290291],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024417732,0.00058871455,0.05333987,0.00088213634,0.00070401444,0.0014540254,0.0053985254,0.6387519,0.014798816,0.0769885,0.007834014,0.19681765],"study_design_scores_gemma":[0.000035574645,0.00022293966,0.0013629022,0.00007114539,0.000098695076,0.0003707916,0.00035258656,0.94803876,0.008408533,0.038853712,0.0021285878,0.000055781235],"about_ca_topic_score_codex":0.0030635858,"about_ca_topic_score_gemma":0.0035129134,"teacher_disagreement_score":0.01830155,"about_ca_system_score_codex":0.001524094,"about_ca_system_score_gemma":0.0017381056,"threshold_uncertainty_score":0.096789},"labels":[],"label_agreement":null},{"id":"W4389519983","doi":"10.18653/v1/2023.findings-emnlp.518","title":"Impact of Co-occurrence on Factual Knowledge of Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Subject (documents); Recall; Object (grammar); Set (abstract data type); Natural language processing; Artificial intelligence; Cognitive psychology; Machine learning; Psychology; World Wide Web; Programming language","score_opus":0.056339912736957626,"score_gpt":0.36570269511717235,"score_spread":0.30936278238021475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389519983","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79931724,0.0058974884,0.16869621,0.0029321352,0.0006172689,0.00027633,0.0034265928,0.012981456,0.0058553093],"genre_scores_gemma":[0.95853996,0.00041028977,0.03483276,0.00060026604,0.00010043706,0.00014670407,0.003537467,0.0008323805,0.0009997505],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9842562,0.007084792,0.0014779639,0.004334016,0.0022603953,0.00058666716],"domain_scores_gemma":[0.8345807,0.12648825,0.0065773274,0.02737236,0.0037644482,0.0012168051],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016914302,0.0014041035,0.0011786973,0.0012307534,0.0010922743,0.0029519224,0.001726368,0.0013857062,0.0017915778],"category_scores_gemma":[0.15019913,0.00089715986,0.00095654215,0.0011687801,0.0016085131,0.005973483,0.0031197045,0.0036013676,0.0010586663],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002386763,0.0005681161,0.25921163,0.0019355952,0.001454416,0.0019467666,0.004387186,0.1887935,0.026245683,0.008209191,0.020981424,0.48387975],"study_design_scores_gemma":[0.00019408185,0.00080138753,0.06750428,0.00072993437,0.0007699878,0.002509305,0.002428983,0.79542196,0.054210577,0.050013267,0.02516281,0.0002534416],"about_ca_topic_score_codex":0.0058420985,"about_ca_topic_score_gemma":0.008581834,"teacher_disagreement_score":0.016914302,"about_ca_system_score_codex":0.0012564947,"about_ca_system_score_gemma":0.0014806244,"threshold_uncertainty_score":0.089452505},"labels":[],"label_agreement":null},{"id":"W4389520062","doi":"10.18653/v1/2023.findings-emnlp.499","title":"Reasoning Makes Good Annotators : An Automatic Task-specific Rules Distilling Framework for Low-resource Relation Extraction","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Zhejiang University; National Natural Science Foundation of China; National Key Research and Development Program of China; Ant Group; National Science Foundation","keywords":"Computer science; Task (project management); Relationship extraction; Pipeline (software); Relation (database); Exploit; Set (abstract data type); Artificial intelligence; Resource (disambiguation); Machine learning; Data mining; Labeled data; Natural language processing; Programming language","score_opus":0.028429755429020744,"score_gpt":0.28635737372351605,"score_spread":0.2579276182944953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520062","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006444475,0.0004662947,0.972138,0.0004131556,0.000086312626,0.0002347494,0.0014883052,0.017107269,0.0016213511],"genre_scores_gemma":[0.057623725,0.0002460445,0.9304918,0.00042909614,0.0000802629,0.00025395965,0.0073253256,0.0009183806,0.0026314505],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99621046,0.0009653009,0.00030679876,0.0016122839,0.0007019433,0.00020324335],"domain_scores_gemma":[0.9943739,0.0024276623,0.00037167643,0.0018510754,0.00072894775,0.00024687688],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004490582,0.0020299251,0.0011223963,0.003985579,0.0014923519,0.0025207072,0.0039364374,0.0020025116,0.0047021178],"category_scores_gemma":[0.011889573,0.0010255231,0.0028646598,0.0022451174,0.0011070927,0.005781526,0.0038062134,0.0036691972,0.0050982],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033119658,0.0004805825,0.004490833,0.00063473225,0.0002721386,0.0007548621,0.0016855648,0.025714383,0.029272292,0.028263409,0.053807132,0.8542928],"study_design_scores_gemma":[0.000104285355,0.00012037652,0.0015784712,0.00014729626,0.00020923142,0.0007077175,0.00039439806,0.8068177,0.028656727,0.083783634,0.077355325,0.00012493304],"about_ca_topic_score_codex":0.005351128,"about_ca_topic_score_gemma":0.01782483,"teacher_disagreement_score":0.005351128,"about_ca_system_score_codex":0.0009813224,"about_ca_system_score_gemma":0.0027304606,"threshold_uncertainty_score":0.023748755},"labels":[],"label_agreement":null},{"id":"W4389520096","doi":"10.18653/v1/2023.findings-emnlp.500","title":"Co-training and Co-distillation for Quality Improvement and Compression of Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Margin (machine learning); Benchmark (surveying); Distillation; Inference; Language model; Performance improvement; Machine learning; Artificial intelligence; Performance prediction; Simulation; Engineering","score_opus":0.11131643593903276,"score_gpt":0.3746026075313697,"score_spread":0.2632861715923369,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520096","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049915433,0.0018811559,0.9317831,0.0008447697,0.00026472582,0.00013686842,0.00045379525,0.011564508,0.0031557376],"genre_scores_gemma":[0.6003213,0.0007212887,0.3891706,0.00077591406,0.00020190804,0.0002952732,0.002575408,0.0010564525,0.004881822],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978421,0.00074222824,0.00013354383,0.00064389914,0.00041250614,0.00022570461],"domain_scores_gemma":[0.9953068,0.002619538,0.0002069777,0.0012316168,0.00047781548,0.00015735555],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031243167,0.0020687643,0.0016449906,0.0013173467,0.00088502624,0.0016474514,0.002956016,0.001786893,0.003320401],"category_scores_gemma":[0.012966932,0.0008130927,0.0014865027,0.0015803537,0.0014359971,0.0043947645,0.0036248132,0.005183468,0.0017705281],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065751077,0.00037372412,0.0025868765,0.00041586606,0.0003002309,0.00032319507,0.0004130351,0.36210132,0.019551825,0.015623977,0.012061898,0.5855906],"study_design_scores_gemma":[0.000030948708,0.00006824964,0.00022633349,0.000017190101,0.000030145236,0.00006220272,0.000028167751,0.9812882,0.009033582,0.007297598,0.0018957441,0.000021637225],"about_ca_topic_score_codex":0.009330746,"about_ca_topic_score_gemma":0.014879549,"teacher_disagreement_score":0.009330746,"about_ca_system_score_codex":0.0012221733,"about_ca_system_score_gemma":0.002469998,"threshold_uncertainty_score":0.0185529},"labels":[],"label_agreement":null},{"id":"W4389520102","doi":"10.18653/v1/2023.findings-emnlp.381","title":"Prompt-Based Editing for Text Style Transfer","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; University of Alberta","funders":"Alliance de recherche numérique du Canada; Alberta Innovates; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Benchmark (surveying); Word (group theory); Style (visual arts); Artificial intelligence; Natural language processing; Autoregressive model; Task (project management); Text generation; Language model; Process (computing); Speech recognition; Linguistics; Programming language","score_opus":0.04103462738800765,"score_gpt":0.2667797799816062,"score_spread":0.22574515259359854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520102","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03260299,0.0003298012,0.9438409,0.00017605793,0.00019249007,0.00021230787,0.00037874744,0.019979618,0.0022870973],"genre_scores_gemma":[0.5267798,0.00030209025,0.46231014,0.00030621546,0.00026305983,0.0003278549,0.0016324243,0.0012783423,0.006800095],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99836534,0.00061304256,0.00011144283,0.00050202676,0.00033423046,0.00007399964],"domain_scores_gemma":[0.99471843,0.0023832542,0.00037964893,0.001465987,0.00083656097,0.00021600514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022760092,0.0014397886,0.000909604,0.0008279162,0.00038122505,0.0011682165,0.0014531526,0.0010278637,0.004735663],"category_scores_gemma":[0.011040325,0.00030324707,0.0007772719,0.00073171983,0.00050706253,0.0022592968,0.0013154235,0.0014600015,0.0025480834],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000676546,0.0003947118,0.002651618,0.00043490593,0.00008926615,0.0003617162,0.0006390101,0.046382774,0.09780301,0.009007646,0.014399196,0.8271595],"study_design_scores_gemma":[0.0001357232,0.0004349484,0.0017105042,0.00002632625,0.000067373156,0.00037905553,0.00015643529,0.8869949,0.06960648,0.025753098,0.014662769,0.00007230443],"about_ca_topic_score_codex":0.0005310243,"about_ca_topic_score_gemma":0.00086021237,"teacher_disagreement_score":0.004735663,"about_ca_system_score_codex":0.00039857195,"about_ca_system_score_gemma":0.00065593293,"threshold_uncertainty_score":0.015842319},"labels":[],"label_agreement":null},{"id":"W4389520117","doi":"10.18653/v1/2023.findings-emnlp.6","title":"Time-Aware Representation Learning for Time-Sensitive Question Answering","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Context (archaeology); Task (project management); Question answering; Sentence; Metric (unit); Baseline (sea); Artificial intelligence; Natural language processing; Language model; Representation (politics); Context model; Code (set theory); Information retrieval; Machine learning; Programming language","score_opus":0.027313516067080882,"score_gpt":0.2972946109753773,"score_spread":0.2699810949082964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520117","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062302936,0.0024969168,0.9124359,0.0010350177,0.00029760844,0.00025730356,0.0035609882,0.014820529,0.0027928068],"genre_scores_gemma":[0.6158236,0.0011480256,0.3544211,0.00079575385,0.000443984,0.00066609925,0.020504687,0.0005716378,0.005625015],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987178,0.00039644286,0.00008851118,0.0005112317,0.00016143137,0.00012470868],"domain_scores_gemma":[0.9971967,0.0016999021,0.00017244623,0.000479329,0.00034061258,0.00011102887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017508976,0.0012031215,0.00088224316,0.0016691213,0.00040218837,0.0011100121,0.0018919036,0.0017254625,0.0037757382],"category_scores_gemma":[0.0078102187,0.0003796541,0.0015726762,0.0015210413,0.00037935382,0.003518428,0.0016426765,0.0024231432,0.0023281777],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048265923,0.0007071536,0.004167639,0.00048716596,0.00016375596,0.00021343064,0.00071143813,0.094155736,0.021418814,0.011584308,0.040425397,0.82548237],"study_design_scores_gemma":[0.00004713421,0.00014472022,0.0010032946,0.000036218065,0.000055531207,0.0001140382,0.00012127628,0.9559314,0.0061215227,0.029009879,0.007387558,0.000027381253],"about_ca_topic_score_codex":0.004668177,"about_ca_topic_score_gemma":0.0052785315,"teacher_disagreement_score":0.004668177,"about_ca_system_score_codex":0.0011798751,"about_ca_system_score_gemma":0.0011088004,"threshold_uncertainty_score":0.012631118},"labels":[],"label_agreement":null},{"id":"W4389520137","doi":"10.18653/v1/2023.findings-emnlp.377","title":"Disentangling Structure and Style: Political Bias Detection in News by Inducing Document Hierarchy","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology; Iran Telecommunication Research Center; National Research Foundation of Korea; National Research Foundation","keywords":"Computer science; Rhetorical question; Generalizability theory; Robustness (evolution); Style (visual arts); Overfitting; Writing style; Artificial intelligence; Journalism; Natural language processing; Sentence; Topic model; Information retrieval; Politics; Hierarchy; Linguistics; Sociology; Psychology; Political science; Literature; Art","score_opus":0.02283336190972059,"score_gpt":0.2674290424661955,"score_spread":0.24459568055647488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520137","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53257346,0.0021626472,0.45302764,0.000670558,0.00012468756,0.00019382175,0.0007802165,0.0015210016,0.008946044],"genre_scores_gemma":[0.91564566,0.00031533255,0.08064757,0.00011907207,0.00022042794,0.00008106848,0.0009020091,0.0001370996,0.001931914],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998965,0.0003488561,0.00006944816,0.00023903161,0.00025739562,0.00012031386],"domain_scores_gemma":[0.99336064,0.0037413386,0.0009919362,0.0006047507,0.0010075376,0.00029376437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002094465,0.0005402716,0.00047520114,0.002838208,0.00065724546,0.0013287821,0.000555197,0.0005988911,0.0009262268],"category_scores_gemma":[0.008309332,0.00027167966,0.00046049771,0.0019527773,0.00044891282,0.0015993249,0.0010675386,0.0009746314,0.00054582494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085664284,0.00044704555,0.1285588,0.0006710392,0.00029815405,0.00039676926,0.0039810897,0.0169827,0.06340415,0.01349572,0.0076570874,0.76325077],"study_design_scores_gemma":[0.00014941754,0.000514611,0.10667228,0.00018225024,0.00040991718,0.0005281966,0.00091693935,0.78160256,0.049209043,0.044127557,0.015568656,0.00011861274],"about_ca_topic_score_codex":0.003085151,"about_ca_topic_score_gemma":0.0076477844,"teacher_disagreement_score":0.003085151,"about_ca_system_score_codex":0.0005469892,"about_ca_system_score_gemma":0.0009396094,"threshold_uncertainty_score":0.011076689},"labels":[],"label_agreement":null},{"id":"W4389520156","doi":"10.18653/v1/2023.findings-emnlp.85","title":"MoqaGPT : Zero-Shot Multi-modal Open-domain Question Answering with Large Language Model","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"Canadian Institute for Advanced Research; Nvidia","keywords":"Computer science; Baseline (sea); Modality (human–computer interaction); Ranking (information retrieval); Task (project management); Modalities; Artificial intelligence; Codebase; Modal; Benchmark (surveying); Shot (pellet); Language model; Information retrieval; Domain (mathematical analysis); Machine learning; Natural language processing; Programming language; Source code; Mathematics","score_opus":0.04094407819874246,"score_gpt":0.31381736393129295,"score_spread":0.2728732857325505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520156","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009559362,0.001906872,0.81686616,0.001623712,0.000500313,0.00094014924,0.014605459,0.14813785,0.0058600926],"genre_scores_gemma":[0.13543421,0.0006705797,0.78263324,0.0020690006,0.00032059537,0.0016143356,0.06424981,0.0036616006,0.009346569],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99698454,0.0010272395,0.0001728915,0.0011217941,0.00051631173,0.00017725298],"domain_scores_gemma":[0.99611723,0.0019952971,0.00012677144,0.0010137813,0.00051108573,0.0002357412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032070195,0.0028580872,0.0016353833,0.002149075,0.0011591226,0.0028798115,0.0054458412,0.004234976,0.022418251],"category_scores_gemma":[0.014897048,0.00095387705,0.002918331,0.0015777661,0.0012433319,0.006410411,0.0070397155,0.004931134,0.012636313],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006013827,0.00071095285,0.0022140613,0.0023924878,0.0005098595,0.00047255613,0.0008841904,0.033454638,0.015533239,0.019497292,0.35488048,0.5688489],"study_design_scores_gemma":[0.00027485297,0.0002807919,0.0011860334,0.00015188887,0.00011706065,0.0006038816,0.00046797536,0.81875974,0.012398405,0.076566935,0.08906853,0.00012381625],"about_ca_topic_score_codex":0.011468716,"about_ca_topic_score_gemma":0.022306904,"teacher_disagreement_score":0.022418251,"about_ca_system_score_codex":0.0018116971,"about_ca_system_score_gemma":0.0026887883,"threshold_uncertainty_score":0.07499653},"labels":[],"label_agreement":null},{"id":"W4389520158","doi":"10.18653/v1/2023.findings-emnlp.404","title":"NASH: A Simple Unified Framework of Structured Pruning for Accelerating Encoder-Decoder Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Encoder; Speedup; Inference; Pruning; Language model; Algorithm; Artificial intelligence; Parallel computing","score_opus":0.06707218725630042,"score_gpt":0.3174647955264347,"score_spread":0.25039260827013426,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520158","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053765164,0.0001996622,0.9908054,0.00013940755,0.000041380892,0.00006479172,0.00009320509,0.0019140582,0.0013655877],"genre_scores_gemma":[0.20542347,0.0003902711,0.78861046,0.0002688452,0.0000974811,0.00030340147,0.0005421262,0.00075977895,0.0036042116],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99874157,0.0004752915,0.000087977045,0.00023788305,0.0003546026,0.000102689846],"domain_scores_gemma":[0.99730396,0.0015868985,0.00013560057,0.00049617566,0.00037286658,0.000104496205],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023752418,0.0011989886,0.00101643,0.0011901933,0.00074423285,0.0015039014,0.0028978111,0.0013821529,0.003999636],"category_scores_gemma":[0.01046352,0.00072064536,0.0010628429,0.0009177654,0.001029797,0.0035739867,0.0025308575,0.0023763776,0.0015035622],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032296948,0.00016405481,0.0015609439,0.00033166434,0.0001634333,0.00041566615,0.0004857667,0.44552904,0.015713505,0.16050497,0.008290324,0.36651763],"study_design_scores_gemma":[0.000021725815,0.000033319382,0.00007410749,0.000014867012,0.000029996067,0.000058505168,0.000019649888,0.9558577,0.0032645727,0.03840536,0.0022090047,0.000011274506],"about_ca_topic_score_codex":0.0052228873,"about_ca_topic_score_gemma":0.014023782,"teacher_disagreement_score":0.0052228873,"about_ca_system_score_codex":0.0011391444,"about_ca_system_score_gemma":0.0023892007,"threshold_uncertainty_score":0.01338017},"labels":[],"label_agreement":null},{"id":"W4389520187","doi":"10.18653/v1/2023.findings-emnlp.413","title":"Responsible AI Considerations in Text Summarization Research: A Review of Current Practices","year":2023,"lang":"en","type":"review","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); Canadian Institute for Advanced Research; McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Microsoft Research","keywords":"Automatic summarization; Work (physics); Computer science; Task (project management); Data science; Focus (optics); Engineering ethics; Artificial intelligence; Engineering","score_opus":0.709494718523932,"score_gpt":0.5701332435127068,"score_spread":0.1393614750112252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520187","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00019502178,0.9943163,0.00091578765,0.0028778324,0.00029809732,0.000062518746,0.00005590399,0.000027795173,0.0012508128],"genre_scores_gemma":[0.0025566686,0.992906,0.0027771967,0.0009642116,0.00026895548,0.00016030976,0.00009348343,0.000029789417,0.0002434433],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9771807,0.009975748,0.005982875,0.0020057205,0.004440656,0.00041425088],"domain_scores_gemma":[0.74694103,0.21023586,0.011893341,0.0043429085,0.024963629,0.0016231731],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.043316353,0.0012344533,0.0028546862,0.021101039,0.0015406786,0.0069089006,0.0033133044,0.003456119,0.0046267435],"category_scores_gemma":[0.13004002,0.0016531959,0.0020227057,0.024448652,0.0038527073,0.009889267,0.003455784,0.0034005742,0.0025629643],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000092544615,0.00005215807,0.00053366227,0.16416678,0.00038954045,0.000109074164,0.002009573,0.00031437675,0.0005016932,0.009424636,0.022210391,0.8001955],"study_design_scores_gemma":[0.000042653177,0.00012510084,0.0016280143,0.26025105,0.0011414142,0.00043791588,0.0016947213,0.00021388447,0.0007216385,0.008399378,0.72526526,0.000078962505],"about_ca_topic_score_codex":0.0071721743,"about_ca_topic_score_gemma":0.013455806,"teacher_disagreement_score":0.95668364,"about_ca_system_score_codex":0.0058447192,"about_ca_system_score_gemma":0.018140359,"threshold_uncertainty_score":0.22908151},"labels":[],"label_agreement":null},{"id":"W4389520208","doi":"10.18653/v1/2023.findings-emnlp.383","title":"NERvous About My Health: Constructing a Bengali Medical Named Entity Recognition Dataset","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Bengali; Computer science; Flexibility (engineering); Natural language processing; Artificial intelligence; Task (project management); Variety (cybernetics); Named-entity recognition; Sentence; F1 score; Data science","score_opus":0.07622441534822924,"score_gpt":0.32647196998203976,"score_spread":0.25024755463381054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520208","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2533259,0.004351308,0.012342507,0.004169881,0.000810576,0.0019055435,0.70000464,0.009852577,0.013237021],"genre_scores_gemma":[0.10347371,0.0008212282,0.013271115,0.0004970234,0.00012151924,0.000824172,0.8767943,0.00019259371,0.004004353],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983064,0.00033154504,0.00026998261,0.00053145946,0.00030949112,0.00025106032],"domain_scores_gemma":[0.9983589,0.0004175179,0.00016194703,0.0004500374,0.00041749625,0.00019406974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009830344,0.001490659,0.00091688306,0.004077439,0.0014463237,0.0013495351,0.0030617544,0.0019712276,0.0039376193],"category_scores_gemma":[0.0037278985,0.00030737743,0.0011146866,0.0038028515,0.0009099767,0.001076214,0.002157197,0.0013833967,0.0063405884],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028205505,0.0020777492,0.069698304,0.0047932714,0.0007652945,0.01038565,0.0035379487,0.017122043,0.030907705,0.0047013145,0.5867975,0.26639283],"study_design_scores_gemma":[0.00044474754,0.0007902062,0.21314578,0.00062713295,0.0004998891,0.0058235624,0.0076463954,0.052452732,0.034079406,0.0034103359,0.680641,0.00043883093],"about_ca_topic_score_codex":0.048837073,"about_ca_topic_score_gemma":0.05434429,"teacher_disagreement_score":0.048837073,"about_ca_system_score_codex":0.0027642103,"about_ca_system_score_gemma":0.0020241868,"threshold_uncertainty_score":0.09710562},"labels":[],"label_agreement":null},{"id":"W4389520242","doi":"10.18653/v1/2023.findings-emnlp.670","title":"LLM aided semi-supervision for efficient Extractive Dialog Summarization","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Waterloo","funders":"","keywords":"Automatic summarization; Computer science; Dialog box; Set (abstract data type); Frame (networking); Natural language processing; Artificial intelligence; Training set; Data set; Information retrieval; Labeled data; World Wide Web; Programming language","score_opus":0.035390883151155586,"score_gpt":0.27779393052864154,"score_spread":0.24240304737748597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520242","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016177213,0.00083558145,0.9460604,0.0003811934,0.00017946205,0.00020198117,0.0027408975,0.032094408,0.0013289289],"genre_scores_gemma":[0.2732201,0.00032085562,0.6941089,0.0004946935,0.0003158363,0.0009319172,0.022857632,0.0014019625,0.0063481606],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971644,0.001256419,0.00018825114,0.0008105431,0.00040078515,0.00017946628],"domain_scores_gemma":[0.99424237,0.0026577134,0.00039717605,0.0011858686,0.0013089243,0.0002078615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028357545,0.0023863977,0.0016277651,0.0018751792,0.0009994383,0.0014167622,0.0022403607,0.0015637043,0.004066616],"category_scores_gemma":[0.010077702,0.0007647925,0.001419472,0.0013224222,0.00066102366,0.0028714887,0.002370397,0.0029935248,0.004994532],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010168618,0.00050298433,0.002491446,0.00090175046,0.00022967787,0.00030337615,0.0011578358,0.07820518,0.06504767,0.0044278945,0.05086284,0.79485244],"study_design_scores_gemma":[0.00008126003,0.00020261788,0.0009919666,0.00004458108,0.0000599594,0.0000804906,0.00023732141,0.94787955,0.029177291,0.008200557,0.012990116,0.000054287288],"about_ca_topic_score_codex":0.0063735186,"about_ca_topic_score_gemma":0.016969433,"teacher_disagreement_score":0.0063735186,"about_ca_system_score_codex":0.001092855,"about_ca_system_score_gemma":0.0023978245,"threshold_uncertainty_score":0.014997065},"labels":[],"label_agreement":null},{"id":"W4389520250","doi":"10.18653/v1/2023.findings-emnlp.753","title":"A Zero-Shot Language Agent for Computer Control with Structured Reflection","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Task (project management); Executable; TRACE (psycholinguistics); Reflection (computer programming); Control (management); Artificial intelligence; Shot (pellet); Human–computer interaction; Programming language","score_opus":0.03325824647450009,"score_gpt":0.28897376028585714,"score_spread":0.25571551381135704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520250","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024332732,0.00009162274,0.9628391,0.0004452897,0.00007711362,0.0001448453,0.00005693653,0.005538534,0.0064738905],"genre_scores_gemma":[0.47250634,0.000089497444,0.5161613,0.00032485017,0.000033888016,0.0003192312,0.00016792153,0.0004235267,0.009973437],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936754,0.00022001404,0.000030659492,0.00017838499,0.00013933117,0.000064059735],"domain_scores_gemma":[0.9989189,0.00046300117,0.00009014156,0.00026763213,0.00011567056,0.00014470381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010557183,0.0006807606,0.0004967694,0.0003147961,0.0005900285,0.0012584794,0.0020859404,0.001284587,0.005199148],"category_scores_gemma":[0.0040532,0.0004120857,0.0006298153,0.0001788084,0.0017200633,0.0017985153,0.0023961323,0.0017370816,0.0010627117],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001051124,0.00082314305,0.003149064,0.00043434044,0.0001653591,0.00081591395,0.003069919,0.37007582,0.05557035,0.25923884,0.014774768,0.29083124],"study_design_scores_gemma":[0.000077377736,0.0001320599,0.000107076594,0.000021637607,0.000022430286,0.000086203625,0.00009830202,0.9438067,0.006547322,0.040157624,0.008914926,0.000028340508],"about_ca_topic_score_codex":0.0025728582,"about_ca_topic_score_gemma":0.0036354526,"teacher_disagreement_score":0.005199148,"about_ca_system_score_codex":0.0005732313,"about_ca_system_score_gemma":0.0015613864,"threshold_uncertainty_score":0.017392814},"labels":[],"label_agreement":null},{"id":"W4389520276","doi":"10.18653/v1/2023.findings-emnlp.703","title":"Investigating the Effect of Pre-finetuning BERT Models on NLI Involving Presuppositions","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Institut de Valorisation des Données; Canadian Institute for Advanced Research","keywords":"Computer science; Leverage (statistics); Presupposition; Artificial intelligence; Exploit; Machine learning; Task (project management); Theoretical computer science","score_opus":0.03633354126024544,"score_gpt":0.27223077294286224,"score_spread":0.2358972316826168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520276","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85880226,0.0028956267,0.1180536,0.0018836103,0.0005066679,0.00042205877,0.0007713972,0.0076178657,0.009046862],"genre_scores_gemma":[0.9541312,0.00029311937,0.041397527,0.00049625366,0.000066459455,0.0001830364,0.0014187925,0.0004258871,0.0015876809],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99817276,0.0007944996,0.00012383897,0.000595796,0.00011118441,0.00020198591],"domain_scores_gemma":[0.97622865,0.018472243,0.00083143223,0.0029822025,0.00086313626,0.00062238454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005622763,0.0028173847,0.0013749957,0.00081574905,0.0006928685,0.002371533,0.0024129695,0.0025434897,0.0034340464],"category_scores_gemma":[0.036430005,0.0007538376,0.0010168867,0.00058000354,0.0010991817,0.004943105,0.0027333822,0.0065576322,0.0018819412],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028083439,0.0028358668,0.03956182,0.0013763201,0.00082993106,0.00063290016,0.0018651176,0.47428834,0.043923322,0.0033481438,0.007674484,0.4208554],"study_design_scores_gemma":[0.00015531613,0.0010242022,0.0044584414,0.00013050322,0.0001880132,0.0001762097,0.00058734376,0.9534301,0.029310102,0.007796358,0.0026596354,0.00008369078],"about_ca_topic_score_codex":0.006225663,"about_ca_topic_score_gemma":0.008959459,"teacher_disagreement_score":0.006225663,"about_ca_system_score_codex":0.0011003787,"about_ca_system_score_gemma":0.0013157662,"threshold_uncertainty_score":0.02973634},"labels":[],"label_agreement":null},{"id":"W4389520290","doi":"10.18653/v1/2023.findings-emnlp.686","title":"Can Large Language Models Fix Data Annotation Errors? An Empirical Study Using Debatepedia for Query-Focused Text Summarization","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Bank of Canada; York University","funders":"Natural Sciences and Engineering Research Council of Canada; York University; Compute Canada","keywords":"Automatic summarization; Computer science; Information retrieval; Relevance (law); Query expansion; RDF query language; Web search query; Task (project management); Query language; Multi-document summarization; Sampling (signal processing); Natural language processing; Web query classification; Language model; Artificial intelligence; Search engine","score_opus":0.15478084264944392,"score_gpt":0.38446822318144025,"score_spread":0.22968738053199633,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520290","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85969025,0.0061280816,0.09926151,0.0036172562,0.0003944303,0.0008652581,0.019497829,0.006594177,0.0039512403],"genre_scores_gemma":[0.8742853,0.00070973916,0.07299417,0.0007225241,0.00022698194,0.0007655544,0.04839371,0.00047027878,0.0014317983],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98216474,0.013152124,0.001029575,0.0021773914,0.0011522416,0.00032386387],"domain_scores_gemma":[0.8639822,0.115533315,0.004398963,0.010845959,0.004471158,0.00076827337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01916592,0.0010435942,0.0010786128,0.002421211,0.0012018959,0.002033363,0.002038692,0.0015118207,0.0016857481],"category_scores_gemma":[0.10646226,0.00040336003,0.0011900465,0.0032940344,0.0012802606,0.004894752,0.00137861,0.0032474513,0.0017294104],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005066824,0.0029634451,0.17486529,0.0060085845,0.0015602766,0.001222858,0.009757915,0.09201641,0.019999832,0.010013716,0.10848893,0.568036],"study_design_scores_gemma":[0.001137693,0.0023867942,0.08164338,0.0005986855,0.0010485327,0.0015770014,0.0070649954,0.7729041,0.031907752,0.017917022,0.08149057,0.00032349004],"about_ca_topic_score_codex":0.00536709,"about_ca_topic_score_gemma":0.0063276896,"teacher_disagreement_score":0.01916592,"about_ca_system_score_codex":0.0013209903,"about_ca_system_score_gemma":0.0013612385,"threshold_uncertainty_score":0.10136026},"labels":[],"label_agreement":null},{"id":"W4389520397","doi":"10.18653/v1/2023.emnlp-main.593","title":"EpiK-Eval: Evaluation for Language Models as Epistemic Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds de recherche du Québec – Nature et technologies; Alliance de recherche numérique du Canada","keywords":"Benchmark (surveying); Computer science; Narrative; Representation (politics); Artificial intelligence; Consolidation (business); Data science; Political science; Linguistics; Business; Geography","score_opus":0.12510718116200836,"score_gpt":0.35659664637991295,"score_spread":0.2314894652179046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520397","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38371617,0.0101818945,0.32138592,0.0056783515,0.0024752081,0.0026912466,0.056447428,0.13878362,0.07864018],"genre_scores_gemma":[0.60375285,0.0013415877,0.30734175,0.0012863749,0.0001838728,0.0015026502,0.07087964,0.0057178247,0.0079934755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99034345,0.005548709,0.00080465095,0.0013761288,0.0015364995,0.0003904911],"domain_scores_gemma":[0.96360993,0.028524354,0.00065978756,0.0038966115,0.0025777172,0.0007315591],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008592129,0.0024417078,0.0010237703,0.002165901,0.00082752696,0.0036319555,0.0051840586,0.0032049592,0.0111661265],"category_scores_gemma":[0.06553539,0.0006397436,0.0012756736,0.0013693171,0.0012053405,0.0063275634,0.003970018,0.004037552,0.0039471933],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004786835,0.0023974257,0.01433416,0.0070702815,0.0011163863,0.00046530474,0.0015679217,0.3469632,0.008318034,0.028726894,0.16064686,0.4236067],"study_design_scores_gemma":[0.00050603354,0.0007374962,0.0022015353,0.00040999,0.00014084022,0.00028341182,0.000740041,0.92406523,0.011170275,0.023172017,0.036478754,0.00009427921],"about_ca_topic_score_codex":0.010493319,"about_ca_topic_score_gemma":0.015044134,"teacher_disagreement_score":0.0111661265,"about_ca_system_score_codex":0.0020967876,"about_ca_system_score_gemma":0.0023252168,"threshold_uncertainty_score":0.045440078},"labels":[],"label_agreement":null},{"id":"W4389520434","doi":"10.18653/v1/2023.findings-emnlp.567","title":"Contrastive Deterministic Autoencoders For Language Modeling","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Waterloo","funders":"","keywords":"Computer science; Language model; Artificial intelligence; Transformer; Mixture model; Machine learning; Representation (politics); Natural language processing","score_opus":0.0410131605871742,"score_gpt":0.29895284353631685,"score_spread":0.25793968294914266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520434","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006368266,0.00041069763,0.9902866,0.00026182397,0.000040945182,0.000015048805,0.00012508893,0.0005950877,0.0018964043],"genre_scores_gemma":[0.573767,0.0011195831,0.41282818,0.00041225262,0.00017058356,0.00020028782,0.00086743355,0.0005528141,0.010081914],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99958235,0.00017725967,0.00002134278,0.00010129537,0.00008474093,0.000033061515],"domain_scores_gemma":[0.9983919,0.0011939087,0.00007635955,0.0001814907,0.00012219376,0.0000340629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010852008,0.0007026681,0.0005115901,0.00041954187,0.00033879754,0.00095458794,0.0011126486,0.00086562295,0.0036404205],"category_scores_gemma":[0.00503181,0.0005333641,0.00091213174,0.00052758853,0.0007495468,0.0016846892,0.0010160499,0.0025782418,0.001320699],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006292897,0.000037799913,0.0005168185,0.00010847268,0.000067613575,0.00007688506,0.0001070621,0.808593,0.0054390933,0.10944639,0.0030381817,0.07250577],"study_design_scores_gemma":[0.0000020080831,0.000006949197,0.000047483616,0.000005689056,0.000003571611,0.000012290997,0.0000033692108,0.97717893,0.0007347897,0.021167217,0.0008335125,0.000004173614],"about_ca_topic_score_codex":0.0041143484,"about_ca_topic_score_gemma":0.007569149,"teacher_disagreement_score":0.0041143484,"about_ca_system_score_codex":0.0010183743,"about_ca_system_score_gemma":0.0007689331,"threshold_uncertainty_score":0.012178421},"labels":[],"label_agreement":null},{"id":"W4389520485","doi":"10.18653/v1/2023.findings-emnlp.652","title":"NEWTON: Are Large Language Models Capable of Physical Reasoning?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Office of Naval Research; Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; National Science Foundation","keywords":"Benchmark (surveying); Computer science; Qualitative reasoning; Consistency (knowledge bases); Automated reasoning; Artificial intelligence; Mainstream; Natural language processing; Data science","score_opus":0.020603424736321544,"score_gpt":0.26548124642573323,"score_spread":0.2448778216894117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520485","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12813851,0.0035666856,0.7582056,0.007859998,0.0004166938,0.0007715374,0.018797034,0.046939824,0.035304103],"genre_scores_gemma":[0.58540785,0.0014783242,0.36607835,0.0016952523,0.0001526764,0.001089279,0.034873843,0.004120563,0.0051038046],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9899232,0.005856794,0.00057843083,0.0014405609,0.0019116984,0.00028922464],"domain_scores_gemma":[0.95736945,0.031737726,0.0014162468,0.00634485,0.0025219652,0.00060977344],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00925414,0.001242631,0.00078542787,0.0017590299,0.0007441027,0.0051373243,0.002785238,0.002105539,0.011338493],"category_scores_gemma":[0.07197121,0.0008942854,0.0013636728,0.00129205,0.0018767767,0.013756449,0.0038209155,0.0026450146,0.004024735],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017104885,0.00066418474,0.023548849,0.0046509155,0.00057811406,0.0007523407,0.008691358,0.09741598,0.0124388235,0.22128348,0.12683539,0.50143003],"study_design_scores_gemma":[0.0003575528,0.00036262829,0.006869333,0.0010150899,0.0002503609,0.0008562351,0.0023278478,0.4781929,0.011187939,0.3506242,0.14774223,0.0002136527],"about_ca_topic_score_codex":0.008215504,"about_ca_topic_score_gemma":0.00941435,"teacher_disagreement_score":0.011338493,"about_ca_system_score_codex":0.0017831654,"about_ca_system_score_gemma":0.0026557066,"threshold_uncertainty_score":0.048941195},"labels":[],"label_agreement":null},{"id":"W4389520668","doi":"10.18653/v1/2023.emnlp-main.404","title":"Efficient Classification of Long Documents via State-Space Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Computation; Robustness (evolution); Artificial intelligence; Machine learning; Transformer; Binary classification; Extrapolation; Binary number; Data mining; Pattern recognition (psychology); Algorithm; Support vector machine; Mathematics","score_opus":0.043263523656582245,"score_gpt":0.2815403646310278,"score_spread":0.23827684097444551,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520668","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045733653,0.00076762284,0.9474617,0.00050009374,0.0000750508,0.00005996665,0.00042194058,0.0034948895,0.0014850885],"genre_scores_gemma":[0.8039085,0.0007591208,0.18475352,0.0003035406,0.0001859322,0.00020879396,0.0023187033,0.00026398688,0.0072978875],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99956924,0.00012064695,0.0000308808,0.00014604465,0.00007872254,0.000054474465],"domain_scores_gemma":[0.9981382,0.0012552672,0.0001296658,0.00020342914,0.00022627231,0.000047070076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009780504,0.00080542674,0.0009607857,0.0010005278,0.00036078363,0.0014003895,0.0011572135,0.0010690362,0.002322664],"category_scores_gemma":[0.003344418,0.0003764559,0.0010806301,0.0011581243,0.00041205742,0.0027048513,0.0007829533,0.0019328411,0.0015390143],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003335788,0.00020257184,0.0020907682,0.00017223277,0.00010831138,0.00011054774,0.00018237212,0.5311338,0.0067314054,0.017913159,0.0065456503,0.43447566],"study_design_scores_gemma":[0.0000031880477,0.00001008613,0.000077071636,0.0000028988416,0.0000047126405,0.000007655113,0.0000062366166,0.9942492,0.0006534086,0.0047624633,0.00022010341,0.0000029753671],"about_ca_topic_score_codex":0.0052124457,"about_ca_topic_score_gemma":0.0068896287,"teacher_disagreement_score":0.0052124457,"about_ca_system_score_codex":0.00087245816,"about_ca_system_score_gemma":0.0009854378,"threshold_uncertainty_score":0.010364234},"labels":[],"label_agreement":null},{"id":"W4389520674","doi":"10.18653/v1/2023.emnlp-main.433","title":"Unified Low-Resource Sequence Labeling by Sample-Aware Dynamic Sparse Finetuning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Cisco Systems","keywords":"Computer science; Leverage (statistics); Sequence (biology); Context (archaeology); Resource (disambiguation); Sample (material); Generalization; Process (computing); Artificial intelligence; Machine learning; Data mining; Programming language","score_opus":0.041925591101345465,"score_gpt":0.27350071694321815,"score_spread":0.23157512584187268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520674","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054248303,0.0006725941,0.9141783,0.0005559371,0.00019013356,0.0002298717,0.0009265548,0.025518438,0.0034798617],"genre_scores_gemma":[0.4667762,0.00033218233,0.5156765,0.0013866118,0.00015826614,0.00065541203,0.0057532466,0.0027870943,0.0064744554],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904984,0.00022805989,0.00004462414,0.00039045888,0.00017771573,0.00010924468],"domain_scores_gemma":[0.9977787,0.0010040638,0.00013684107,0.0006336314,0.00029405567,0.00015275965],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012130701,0.0019746122,0.001524697,0.00084663916,0.00061025,0.0011798267,0.0028853389,0.0015902618,0.005019992],"category_scores_gemma":[0.007789608,0.0008105885,0.0011190509,0.00088482595,0.0010655737,0.003178071,0.0026855066,0.0033557587,0.0029044915],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062196056,0.00058780337,0.0038968187,0.00036282238,0.00014221248,0.00029994984,0.00033149915,0.414381,0.03396512,0.0067611863,0.020905102,0.51774454],"study_design_scores_gemma":[0.000039380393,0.00007092521,0.00024196757,0.000014310662,0.000017474695,0.000053864616,0.000041578754,0.98537004,0.005542896,0.006462558,0.002127429,0.000017496255],"about_ca_topic_score_codex":0.0068022883,"about_ca_topic_score_gemma":0.017694611,"teacher_disagreement_score":0.0068022883,"about_ca_system_score_codex":0.0010922878,"about_ca_system_score_gemma":0.0019498132,"threshold_uncertainty_score":0.016793549},"labels":[],"label_agreement":null},{"id":"W4389520703","doi":"10.18653/v1/2023.emnlp-main.489","title":"TheoremQA: A Theorem-driven Question Answering Dataset","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Benchmark (surveying); Computer science; Domain (mathematical analysis); Code (set theory); Question answering; Quality (philosophy); Baseline (sea); Artificial intelligence; Programming language; Mathematics; Epistemology","score_opus":0.026374327893978304,"score_gpt":0.28410811597698477,"score_spread":0.25773378808300645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389520703","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05857262,0.0029907909,0.018167129,0.0031689873,0.00059937156,0.0009776126,0.8729107,0.026489131,0.016123718],"genre_scores_gemma":[0.04438863,0.00030746253,0.02535624,0.0008411408,0.00007660884,0.0006344652,0.9247033,0.00044278678,0.0032493072],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973924,0.0007051654,0.00028395726,0.00073846534,0.00072380854,0.0001561198],"domain_scores_gemma":[0.9939016,0.003154127,0.00029940184,0.0010224148,0.0011735813,0.00044883127],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026271334,0.0022673544,0.00067072274,0.0028072458,0.0011181362,0.0016554428,0.003836731,0.0042927773,0.012147965],"category_scores_gemma":[0.0134154875,0.00045599285,0.0021690466,0.002090588,0.00093259825,0.0030199012,0.0022664226,0.0029058508,0.010776609],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061630754,0.0007871939,0.011959096,0.0031498692,0.0002252541,0.0005539394,0.0006478885,0.0133029595,0.0051467465,0.0077479235,0.904152,0.051710833],"study_design_scores_gemma":[0.0011994917,0.0006924663,0.030208793,0.0004893105,0.00015181929,0.0013307092,0.0011209553,0.14890319,0.013502869,0.01871851,0.7834764,0.00020546163],"about_ca_topic_score_codex":0.01962773,"about_ca_topic_score_gemma":0.038968764,"teacher_disagreement_score":0.01962773,"about_ca_system_score_codex":0.0026436625,"about_ca_system_score_gemma":0.0022334282,"threshold_uncertainty_score":0.040638983},"labels":[],"label_agreement":null},{"id":"W4389521016","doi":"10.18653/v1/2023.conll-1.12","title":"On the utility of enhancing BERT syntactic bias with Token Reordering Pretraining","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; University of Waterloo","funders":"","keywords":"Security token; Computer science; Artificial intelligence; Natural language processing; Natural (archaeology); Parsing; Natural language; Computer security; History","score_opus":0.07364391134414623,"score_gpt":0.2643656723281962,"score_spread":0.19072176098404997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389521016","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24782534,0.007316779,0.68612283,0.0052898647,0.00153204,0.00024907882,0.0016540195,0.018157735,0.0318523],"genre_scores_gemma":[0.7810657,0.0014279784,0.1960148,0.0017060105,0.0004971362,0.00018063246,0.0031400593,0.0016625327,0.0143053075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991289,0.00037627947,0.00005246883,0.00021027762,0.00011302613,0.0001190398],"domain_scores_gemma":[0.9938601,0.004594915,0.00012702912,0.0005752637,0.00069714006,0.00014552243],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021856814,0.0013464956,0.0010584363,0.0009275772,0.00081119896,0.0013265914,0.0021581294,0.0015611709,0.007687429],"category_scores_gemma":[0.008187461,0.0005926783,0.00046868404,0.000988168,0.0008267128,0.0041897316,0.0017473582,0.002636477,0.004452662],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016187412,0.00051160285,0.004589772,0.00024422028,0.00015935439,0.00023324805,0.00027606115,0.06530293,0.020825127,0.008581466,0.026346374,0.871311],"study_design_scores_gemma":[0.00020740119,0.00030963458,0.0017808394,0.000081225204,0.00017303867,0.00014151966,0.00022430549,0.9435491,0.020851921,0.026224451,0.0064076236,0.000048915022],"about_ca_topic_score_codex":0.01144298,"about_ca_topic_score_gemma":0.023230912,"teacher_disagreement_score":0.01144298,"about_ca_system_score_codex":0.00067038275,"about_ca_system_score_gemma":0.0015227455,"threshold_uncertainty_score":0.02571696},"labels":[],"label_agreement":null},{"id":"W4389523693","doi":"10.18653/v1/2023.findings-emnlp.549","title":"Open Domain Multi-document Summarization: A Comprehensive Study of Model Brittleness under Retrieval","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Alliance de recherche numérique du Canada","keywords":"Automatic summarization; Computer science; Information retrieval; Domain (mathematical analysis); Multi-document summarization; Task (project management); Set (abstract data type); Natural language processing","score_opus":0.12228804861294826,"score_gpt":0.3515873804412257,"score_spread":0.22929933182827747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523693","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055661093,0.005778699,0.9322712,0.0018284577,0.000113968075,0.00021447949,0.00076229667,0.0009787817,0.0023910778],"genre_scores_gemma":[0.7984326,0.0036421802,0.18718746,0.0007630275,0.00093828526,0.0005091786,0.0029259138,0.0005960769,0.0050052595],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99595994,0.0020461208,0.00026290893,0.0010669371,0.0004571102,0.00020701873],"domain_scores_gemma":[0.94640964,0.042343203,0.0036878153,0.004390661,0.002504072,0.0006645972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011340726,0.0015521632,0.0029797973,0.002636527,0.0009731255,0.0028561144,0.003377906,0.002288523,0.0024565975],"category_scores_gemma":[0.044764955,0.0009957602,0.001741787,0.002630674,0.0024692593,0.007392485,0.0022390478,0.0034576682,0.00066814147],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003185808,0.00016787164,0.0023861201,0.0007021613,0.00036024477,0.00020535056,0.00056694896,0.88256544,0.0018500686,0.0341318,0.0035628898,0.0731825],"study_design_scores_gemma":[0.00001744287,0.00008208469,0.00047745404,0.00003280879,0.000041370302,0.000058229467,0.00005199579,0.96458924,0.00047570968,0.033135615,0.0010155347,0.000022528622],"about_ca_topic_score_codex":0.007681186,"about_ca_topic_score_gemma":0.0044436078,"teacher_disagreement_score":0.011340726,"about_ca_system_score_codex":0.0027408337,"about_ca_system_score_gemma":0.0014698985,"threshold_uncertainty_score":0.05997616},"labels":[],"label_agreement":null},{"id":"W4389523719","doi":"10.18653/v1/2023.findings-emnlp.772","title":"Asking Clarification Questions to Handle Ambiguity in Open-Domain QA","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Seoul National University","keywords":"Ambiguity; Computer science; Ask price; Pipeline (software); Question answering; Process (computing); Domain (mathematical analysis); Interpretation (philosophy); Open domain; Information retrieval; Data science","score_opus":0.07022568823081114,"score_gpt":0.33584743569458186,"score_spread":0.2656217474637707,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523719","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18109189,0.005596637,0.73390317,0.0038602708,0.0007796484,0.0023595246,0.014587511,0.04205531,0.01576609],"genre_scores_gemma":[0.38592625,0.00067098025,0.5694616,0.0014271622,0.000269791,0.0012882114,0.034867804,0.0015149572,0.004573287],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9841955,0.009796237,0.001093128,0.0027970548,0.0015840518,0.00053412415],"domain_scores_gemma":[0.9419946,0.043098655,0.0018985933,0.007899624,0.0039210375,0.0011875927],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011480681,0.0026474923,0.0016796706,0.0036794199,0.0019436512,0.0032118994,0.0031604383,0.0046336777,0.009068153],"category_scores_gemma":[0.060529415,0.00089974597,0.0016636092,0.0020657,0.0016001497,0.008064264,0.006540736,0.0053003607,0.006376452],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021642654,0.0014806984,0.023825899,0.0076430663,0.00028443753,0.0017269796,0.019848488,0.02849014,0.060521685,0.026202748,0.12896013,0.69885147],"study_design_scores_gemma":[0.0007282524,0.0009903774,0.02179827,0.0014522489,0.0002949404,0.0041619595,0.007925543,0.43263093,0.08316026,0.092164055,0.35419175,0.0005014231],"about_ca_topic_score_codex":0.0033214681,"about_ca_topic_score_gemma":0.004336663,"teacher_disagreement_score":0.011480681,"about_ca_system_score_codex":0.0012899725,"about_ca_system_score_gemma":0.0018729683,"threshold_uncertainty_score":0.06071633},"labels":[],"label_agreement":null},{"id":"W4389523790","doi":"10.18653/v1/2023.findings-emnlp.993","title":"Qualitative Code Suggestion: A Human-Centric Approach to Qualitative Coding","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Fonds de recherche du Québec – Nature et technologies; Fonds de Recherche du Québec - Santé; Canadian Institute for Advanced Research; Nvidia","keywords":"Coding (social sciences); Computer science; Qualitative research; Natural language processing; Qualitative analysis; Annotation; Task (project management); Artificial intelligence; Information retrieval; Human–computer interaction","score_opus":0.21009077325139738,"score_gpt":0.4354944176461925,"score_spread":0.22540364439479513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523790","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015800554,0.00020965128,0.9502575,0.00481963,0.000326906,0.01752279,0.0018511033,0.0007110165,0.008500824],"genre_scores_gemma":[0.05747029,0.00023550843,0.90422183,0.0013587205,0.00008785422,0.03145578,0.0010337922,0.00039228596,0.0037439705],"study_design_codex":"qualitative","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.7664721,0.19293568,0.009774592,0.011966826,0.017014796,0.0018360238],"domain_scores_gemma":[0.5960613,0.26791194,0.018584758,0.038641606,0.07573517,0.0030652739],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.17132728,0.0019402825,0.0014748008,0.008824011,0.0069844117,0.0084556835,0.0056412453,0.0021660198,0.008725281],"category_scores_gemma":[0.27201238,0.0013879978,0.0015093962,0.008209231,0.013450193,0.00656619,0.009719705,0.0052319155,0.0029900388],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007736986,0.00040441254,0.0054026027,0.005353568,0.00012529035,0.0007260214,0.4642301,0.004039415,0.021402273,0.21380034,0.031782426,0.25195998],"study_design_scores_gemma":[0.00048240487,0.00053170946,0.0050771087,0.0044233524,0.00011032563,0.0006890159,0.18274419,0.038821712,0.028380698,0.3564818,0.38164195,0.0006157168],"about_ca_topic_score_codex":0.0072514275,"about_ca_topic_score_gemma":0.012020747,"teacher_disagreement_score":0.17132728,"about_ca_system_score_codex":0.01170686,"about_ca_system_score_gemma":0.025536815,"threshold_uncertainty_score":0.9060761},"labels":[],"label_agreement":null},{"id":"W4389523818","doi":"10.18653/v1/2023.findings-emnlp.788","title":"Do “English” Named Entity Recognizers Work Well on Global Englishes?","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Computer science; Artificial intelligence; Natural language processing; Named-entity recognition; Named entity; Transformer; Test (biology); Training set; Test data; Botany","score_opus":0.048995248739012315,"score_gpt":0.2751326913912673,"score_spread":0.22613744265225497,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523818","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68920606,0.011428084,0.17217071,0.006380021,0.0019099648,0.00040862363,0.019099377,0.024850762,0.07454641],"genre_scores_gemma":[0.8765482,0.0030626892,0.05712051,0.0015345674,0.00034734234,0.000110576555,0.04122766,0.002239013,0.017809471],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99727005,0.000734397,0.00016553624,0.0012500499,0.0002624591,0.00031752346],"domain_scores_gemma":[0.9955787,0.0017106791,0.00020444003,0.0015499784,0.0007255117,0.00023056132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005741962,0.0023422274,0.0016044413,0.0019583362,0.00080791616,0.002848593,0.001303523,0.0017155396,0.0052249883],"category_scores_gemma":[0.012597292,0.00056701485,0.0016100277,0.0022002407,0.00081492634,0.0103703,0.0016009392,0.0018475157,0.010058093],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015781999,0.00059056294,0.108420365,0.001562594,0.0018586149,0.0010226193,0.0016158979,0.04080081,0.031666715,0.0039903754,0.08631001,0.72058326],"study_design_scores_gemma":[0.00043480535,0.0024226864,0.17287438,0.0013931735,0.0028871563,0.003200332,0.011321482,0.47006726,0.1033325,0.023122553,0.20833753,0.00060623186],"about_ca_topic_score_codex":0.013932423,"about_ca_topic_score_gemma":0.038338188,"teacher_disagreement_score":0.013932423,"about_ca_system_score_codex":0.00058696617,"about_ca_system_score_gemma":0.0008021921,"threshold_uncertainty_score":0.030366778},"labels":[],"label_agreement":null},{"id":"W4389523824","doi":"10.18653/v1/2023.emnlp-main.797","title":"We are Who We Cite: Bridges of Influence Between Natural Language Processing and Other Academic Fields","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; Institute for Work & Health","funders":"Deutscher Akademischer Austauschdienst; Niedersächsische Ministerium für Wissenschaft und Kultur","keywords":"Citation; Artificial intelligence; Field (mathematics); Computer science; Natural language processing; Computational linguistics; Diversity (politics); Linguistics; Library science; Mathematics; Sociology; Philosophy; Anthropology","score_opus":0.031060447987144506,"score_gpt":0.30383223767134954,"score_spread":0.272771789684205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523824","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8399721,0.030802405,0.010625514,0.023980638,0.0011837716,0.0000638325,0.003427713,0.00041429923,0.0895297],"genre_scores_gemma":[0.9915679,0.0026848665,0.0012949691,0.0006614834,0.000937223,0.000020852765,0.0007466267,0.00006005668,0.0020260175],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99342716,0.0017316663,0.0006785676,0.0012706906,0.0023132015,0.0005786866],"domain_scores_gemma":[0.8613404,0.07701906,0.027630486,0.007936431,0.016658477,0.009415092],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0086960755,0.0003127434,0.00064327253,0.024938062,0.0027774097,0.009704439,0.00072063337,0.0011187152,0.0059739305],"category_scores_gemma":[0.09384006,0.00019004132,0.0004553283,0.029987115,0.0026809794,0.008897337,0.0047301445,0.0011235013,0.0013856771],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014313472,0.00007483515,0.7003681,0.00062388595,0.00049890415,0.00048724434,0.015779015,0.0010705946,0.0013055367,0.04138142,0.018387431,0.2198799],"study_design_scores_gemma":[0.00004215993,0.00013747903,0.7276593,0.00085767126,0.00051120645,0.0010723079,0.013591822,0.0035039103,0.0022936894,0.06523107,0.18495198,0.00014748883],"about_ca_topic_score_codex":0.005776376,"about_ca_topic_score_gemma":0.0058155223,"teacher_disagreement_score":0.97506195,"about_ca_system_score_codex":0.0018367235,"about_ca_system_score_gemma":0.0019977288,"threshold_uncertainty_score":0.04598981},"labels":[],"label_agreement":null},{"id":"W4389523869","doi":"10.18653/v1/2023.conll-babylm.5","title":"Grammar induction pretraining for language modeling in low resource contexts","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Baseline (sea); Computer science; Hyperparameter; Grammar; Context (archaeology); Natural language processing; Language model; Artificial intelligence; Resource (disambiguation); Linguistics","score_opus":0.05666310754084649,"score_gpt":0.28651823205821264,"score_spread":0.22985512451736614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523869","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061842706,0.001066922,0.9045547,0.0012033613,0.0005403841,0.00016570545,0.002190624,0.023467563,0.0049680546],"genre_scores_gemma":[0.5532662,0.00073257904,0.42099714,0.0009307278,0.0002944894,0.00042977068,0.009697177,0.0024814557,0.011170374],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99906904,0.00040522453,0.00003616014,0.00028100715,0.000101677906,0.00010679701],"domain_scores_gemma":[0.9981255,0.001186399,0.000068190064,0.00032790628,0.0002197033,0.0000723872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015525994,0.0013314341,0.0008060817,0.00058361713,0.00048673255,0.0009686638,0.0017403167,0.001084182,0.005702651],"category_scores_gemma":[0.005285725,0.00062135194,0.0011786973,0.00085586106,0.00060967816,0.0032720538,0.001567214,0.004848823,0.0046133255],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006130951,0.0004414192,0.0040452145,0.0005848155,0.00027384696,0.00068535144,0.0006915895,0.25530648,0.0454731,0.023224775,0.067884706,0.6007755],"study_design_scores_gemma":[0.000043684813,0.00013895908,0.0006799122,0.000037752445,0.000053586984,0.00015779288,0.00010922209,0.94898254,0.0142100025,0.02535098,0.010198553,0.000037019705],"about_ca_topic_score_codex":0.0048146234,"about_ca_topic_score_gemma":0.011865313,"teacher_disagreement_score":0.005702651,"about_ca_system_score_codex":0.0007510315,"about_ca_system_score_gemma":0.0013697448,"threshold_uncertainty_score":0.019077241},"labels":[],"label_agreement":null},{"id":"W4389523905","doi":"10.18653/v1/2023.emnlp-main.83","title":"How Does Generative Retrieval Scale to Millions of Passages?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Generative grammar; Transformer; Information retrieval; Search engine indexing; Document retrieval; Artificial intelligence; Encoder; Generative model; Ranking (information retrieval); Task (project management); Natural language processing","score_opus":0.030232360065957312,"score_gpt":0.2627108354550887,"score_spread":0.2324784753891314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523905","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20042937,0.0111095365,0.73249185,0.014755093,0.0009918824,0.00048270603,0.0014593054,0.012764617,0.025515668],"genre_scores_gemma":[0.77762604,0.0046845237,0.20277375,0.0021450277,0.0010023983,0.00036049794,0.0026170206,0.0017829767,0.0070079006],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99634165,0.0014852694,0.00022338124,0.00084705214,0.00083574577,0.0002668724],"domain_scores_gemma":[0.9791266,0.012178774,0.0007178791,0.0057187667,0.0018142408,0.00044391208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005725068,0.0009495816,0.0019126968,0.0019851255,0.0010370654,0.004294151,0.0023118034,0.0021871834,0.004953386],"category_scores_gemma":[0.05261387,0.00092704355,0.0010656716,0.0022593173,0.0021125115,0.014587314,0.0026008822,0.0029537685,0.004444962],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065700023,0.00056689896,0.010791504,0.0014856973,0.00038079053,0.00083091203,0.0025198248,0.14601786,0.035596974,0.067542285,0.0568368,0.67677355],"study_design_scores_gemma":[0.00017306885,0.00039501244,0.0041756458,0.0001613396,0.00022892152,0.0013870454,0.0012753468,0.74316627,0.01948647,0.20064928,0.028753558,0.00014812754],"about_ca_topic_score_codex":0.0042614853,"about_ca_topic_score_gemma":0.003892391,"teacher_disagreement_score":0.005725068,"about_ca_system_score_codex":0.0011465389,"about_ca_system_score_gemma":0.0010515973,"threshold_uncertainty_score":0.030277431},"labels":[],"label_agreement":null},{"id":"W4389523908","doi":"10.18653/v1/2023.emnlp-main.278","title":"Can Large Language Models Capture Dissenting Human Voices?","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Representativeness heuristic; Multinomial distribution; Computer science; Inference; Population; Distribution (mathematics); Scope (computer science); Weighting; Natural language processing; Artificial intelligence; Econometrics; Statistics; Mathematics; Sociology; Demography","score_opus":0.0248647251376716,"score_gpt":0.28294975052646654,"score_spread":0.25808502538879496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523908","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1633363,0.0011216506,0.8270625,0.0019909842,0.00016885511,0.00013722274,0.0004937119,0.0011776717,0.004511023],"genre_scores_gemma":[0.9223078,0.00031830763,0.07338276,0.0007018457,0.0001728654,0.00015510562,0.0010991642,0.00031165086,0.0015504477],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99551845,0.0028107641,0.00011149445,0.001053515,0.0003190666,0.00018663902],"domain_scores_gemma":[0.97166044,0.022676956,0.0012700575,0.0030727498,0.0009010786,0.00041875624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0103013655,0.0010625235,0.0010357196,0.0013277461,0.00080827804,0.0027285225,0.0015996287,0.001691535,0.0022181082],"category_scores_gemma":[0.055934392,0.00069636235,0.00079490605,0.00095308985,0.0016287845,0.0060098246,0.0021558027,0.0030618182,0.0015162547],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014110723,0.00061262486,0.07220762,0.0009151993,0.0010207824,0.0008213814,0.012346293,0.33466202,0.025486048,0.07528764,0.01484749,0.46038172],"study_design_scores_gemma":[0.000046921567,0.0001011347,0.0074718273,0.000075896365,0.000063268686,0.00025963894,0.0009054126,0.89323115,0.002706966,0.08965094,0.005427269,0.000059557875],"about_ca_topic_score_codex":0.002956276,"about_ca_topic_score_gemma":0.0040840236,"teacher_disagreement_score":0.0103013655,"about_ca_system_score_codex":0.00064241025,"about_ca_system_score_gemma":0.0009066452,"threshold_uncertainty_score":0.05447948},"labels":[],"label_agreement":null},{"id":"W4389523920","doi":"10.18653/v1/2023.nllp-1.25","title":"A Comparative Study of Prompting Strategies for Legal Text Classification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; Thomson Reuters (Canada)","funders":"","keywords":"Computer science; Domain (mathematical analysis); Natural language processing; Artificial intelligence; Mathematics","score_opus":0.20016424545340003,"score_gpt":0.3770997134860705,"score_spread":0.1769354680326705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523920","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56246185,0.011894936,0.36024296,0.0026896256,0.0009712367,0.0011551322,0.0023618669,0.049801532,0.008420862],"genre_scores_gemma":[0.7757332,0.0016871677,0.2111253,0.0006070775,0.00022274659,0.0003816806,0.005425218,0.0007541657,0.004063506],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99509615,0.0025603578,0.00035271913,0.001235095,0.0005270462,0.0002285971],"domain_scores_gemma":[0.97633696,0.018393725,0.0007698622,0.0017457692,0.0018338651,0.0009198785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077181063,0.0019645193,0.0010609056,0.002144696,0.0007829757,0.0019405078,0.0016235961,0.0016693496,0.0032977264],"category_scores_gemma":[0.0329393,0.00039612784,0.00080707815,0.0012761974,0.0006230952,0.005454226,0.0013958479,0.002945954,0.0025745425],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026448907,0.0016611345,0.01890971,0.0011795042,0.000210873,0.00023964763,0.0014395912,0.059814483,0.016575536,0.0030364944,0.015219962,0.8790681],"study_design_scores_gemma":[0.00030195678,0.00178609,0.0067484668,0.00015440161,0.0002162844,0.0003439239,0.0016399901,0.9367962,0.029489147,0.008749821,0.013652705,0.00012107812],"about_ca_topic_score_codex":0.003191395,"about_ca_topic_score_gemma":0.005346255,"teacher_disagreement_score":0.0077181063,"about_ca_system_score_codex":0.0013371126,"about_ca_system_score_gemma":0.0021881952,"threshold_uncertainty_score":0.040817738},"labels":[],"label_agreement":null},{"id":"W4389523922","doi":"10.18653/v1/2023.findings-emnlp.928","title":"Learning to love diligent trolls: Accounting for rater effects in the dialogue safety task","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Chatbot; Utterance; Task (project management); Consistency (knowledge bases); Offensive; Artificial intelligence; Human–computer interaction; Natural language processing; Machine learning; Information retrieval; Operations research; Engineering","score_opus":0.018366653743513407,"score_gpt":0.26133052937580664,"score_spread":0.24296387563229324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523922","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45201555,0.0024087965,0.529739,0.0018039151,0.0006372621,0.000513139,0.00096069346,0.003834484,0.008087158],"genre_scores_gemma":[0.9364418,0.00020824424,0.056252003,0.00028330888,0.00030007452,0.00027945865,0.0010238442,0.00039027203,0.004821009],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98603815,0.008389219,0.00066012505,0.003253255,0.0010680443,0.0005912956],"domain_scores_gemma":[0.90661377,0.06899269,0.00633889,0.009923501,0.006155253,0.0019758244],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.027665362,0.0018251167,0.0012095838,0.0014107857,0.0012409306,0.0033435137,0.0019206665,0.0021241766,0.002862794],"category_scores_gemma":[0.10874584,0.0009468411,0.0008488925,0.00095156144,0.0010279509,0.004088002,0.003024969,0.004133919,0.002217761],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024804198,0.0010284844,0.24084295,0.00069961586,0.00095054606,0.0007602585,0.0076833083,0.1377322,0.020300267,0.007250353,0.015234206,0.56503737],"study_design_scores_gemma":[0.00005673113,0.0003403579,0.04085259,0.00008261528,0.00014441817,0.0002503002,0.0005946259,0.9390028,0.0045148814,0.0105748875,0.0034916173,0.000094101786],"about_ca_topic_score_codex":0.008785086,"about_ca_topic_score_gemma":0.0104527725,"teacher_disagreement_score":0.9723346,"about_ca_system_score_codex":0.0011349455,"about_ca_system_score_gemma":0.0013024489,"threshold_uncertainty_score":0.14631015},"labels":[],"label_agreement":null},{"id":"W4389523983","doi":"10.18653/v1/2023.emnlp-main.362","title":"Fast and Robust Early-Exiting Framework for Autoregressive Language Models with Synchronized Parallel Decoding","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Security token; Decoding methods; Autoregressive model; Language model; Inference; Estimator; Algorithm; Artificial intelligence; Mathematics","score_opus":0.03884542175441516,"score_gpt":0.2738687050378683,"score_spread":0.23502328328345315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389523983","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036176022,0.000117361196,0.99427503,0.00006196868,0.000018121636,0.000026523669,0.000040215,0.001603261,0.00023998969],"genre_scores_gemma":[0.275182,0.00027405706,0.71941626,0.00021377423,0.00011737813,0.0002887572,0.000598691,0.0009164712,0.0029925716],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99808645,0.00070137426,0.00013692743,0.00044911506,0.00042042727,0.00020566359],"domain_scores_gemma":[0.99609166,0.0023696418,0.00024523473,0.0005406658,0.00057375635,0.00017901404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004048137,0.0016428885,0.0017695458,0.0009887561,0.0006137304,0.0017066629,0.0034210193,0.0015815625,0.003104839],"category_scores_gemma":[0.011923721,0.0010832452,0.0011229255,0.00084073795,0.0011037961,0.0032410426,0.0027337656,0.0035995303,0.0015413808],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063203374,0.00020317447,0.0016445679,0.0001826247,0.00017327322,0.00027574613,0.00042755285,0.5909691,0.016332356,0.053844426,0.00391916,0.33139595],"study_design_scores_gemma":[0.000013675818,0.00002223915,0.000057656853,0.000004070278,0.000008769007,0.000018785931,0.000006737899,0.9896666,0.001791153,0.008034283,0.00036544973,0.000010534717],"about_ca_topic_score_codex":0.008069312,"about_ca_topic_score_gemma":0.010188453,"teacher_disagreement_score":0.008069312,"about_ca_system_score_codex":0.0012369397,"about_ca_system_score_gemma":0.003243371,"threshold_uncertainty_score":0.021408856},"labels":[],"label_agreement":null},{"id":"W4389524117","doi":"10.18653/v1/2023.conll-babylm.18","title":"McGill BabyLM Shared Task Submission: The Effects of Data Formatting and Structural Biases","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Natural Sciences and Engineering Research Council of Canada; McGill University; Nvidia","keywords":"Disk formatting; Computer science; Task (project management); Track (disk drive); Operating system; Engineering","score_opus":0.07711472960388542,"score_gpt":0.30331774633144254,"score_spread":0.2262030167275571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524117","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46476245,0.0040726587,0.16174494,0.009592332,0.019425353,0.0033514884,0.102228776,0.17754443,0.057277612],"genre_scores_gemma":[0.57220274,0.00041826905,0.1485661,0.004097337,0.0012057459,0.0033049157,0.20424053,0.028663471,0.03730097],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9790395,0.010276059,0.0013503341,0.0037976776,0.0040054913,0.0015309005],"domain_scores_gemma":[0.94986886,0.023418799,0.0012843371,0.012661517,0.0092529,0.0035136966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025943914,0.00364847,0.0021894115,0.0010137376,0.0027277402,0.0042420826,0.0046223802,0.002835024,0.019822156],"category_scores_gemma":[0.082418956,0.0013773515,0.0014213849,0.0018226356,0.0012085169,0.0039285067,0.0062727374,0.005222174,0.015072494],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0067442735,0.0010474115,0.012204349,0.0015680538,0.0007334952,0.00091468357,0.00096547895,0.036625158,0.024009941,0.002644146,0.7297263,0.1828166],"study_design_scores_gemma":[0.004679872,0.004716867,0.040428814,0.0005608105,0.00059049943,0.002020924,0.0018935208,0.47366455,0.081030466,0.02026311,0.369117,0.0010336237],"about_ca_topic_score_codex":0.0253084,"about_ca_topic_score_gemma":0.057194073,"teacher_disagreement_score":0.025943914,"about_ca_system_score_codex":0.002726643,"about_ca_system_score_gemma":0.0042076786,"threshold_uncertainty_score":0.1372062},"labels":[],"label_agreement":null},{"id":"W4389524193","doi":"10.18653/v1/2023.findings-emnlp.796","title":"Using In-Context Learning to Improve Dialogue Safety","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"SAFER; Computer science; Context (archaeology); Human–computer interaction; Artificial intelligence; Machine learning; Deep learning; Risk analysis (engineering); Computer security","score_opus":0.04731199200695141,"score_gpt":0.28403466662608307,"score_spread":0.23672267461913166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524193","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060533263,0.00055274315,0.9232625,0.00061971275,0.0001562966,0.0002084898,0.00014346973,0.01047152,0.004051848],"genre_scores_gemma":[0.74880916,0.00023241121,0.24453035,0.00046161504,0.00014133187,0.00023439694,0.0004862214,0.00065163523,0.0044529084],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9965288,0.001858251,0.00016989037,0.0007280143,0.0005287516,0.00018633006],"domain_scores_gemma":[0.99125284,0.0053768135,0.0005773013,0.0011916247,0.0012298098,0.00037175431],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040860027,0.0018386749,0.0011462179,0.0009193561,0.0010198132,0.0015729136,0.002093447,0.0016474905,0.0050732596],"category_scores_gemma":[0.01807677,0.00047586203,0.0008495783,0.00030084432,0.00093274616,0.0035467884,0.002853321,0.0029678193,0.0017164716],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012318613,0.0013846002,0.0074309316,0.00072571984,0.000225646,0.00034635913,0.002857509,0.23931094,0.060763422,0.009751892,0.0074744574,0.66849667],"study_design_scores_gemma":[0.000057790672,0.0003728815,0.0006217996,0.000042848194,0.00008087186,0.000106546126,0.00029843923,0.9587466,0.023480019,0.011533311,0.0046163807,0.000042512722],"about_ca_topic_score_codex":0.0032504036,"about_ca_topic_score_gemma":0.004166539,"teacher_disagreement_score":0.0050732596,"about_ca_system_score_codex":0.0010191372,"about_ca_system_score_gemma":0.0017012286,"threshold_uncertainty_score":0.021609128},"labels":[],"label_agreement":null},{"id":"W4389524278","doi":"10.18653/v1/2023.findings-emnlp.180","title":"CASE: Commonsense-Augmented Score with an Expanded Answer Space","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Commonsense knowledge; Space (punctuation); Commonsense reasoning; Weighting; Artificial intelligence; Question answering; Measure (data warehouse); Natural language processing; Machine learning; Data mining; Domain knowledge","score_opus":0.06141825289089449,"score_gpt":0.2746300107929852,"score_spread":0.21321175790209074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524278","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07335749,0.0010617662,0.8744979,0.001146455,0.0003857957,0.0005082097,0.0017472734,0.031187164,0.016107874],"genre_scores_gemma":[0.58809716,0.0001974832,0.382264,0.0007180008,0.0004157325,0.0006968992,0.0061489646,0.0014035979,0.020058217],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9950919,0.0020169893,0.00027125914,0.0010009524,0.0013393284,0.00027950652],"domain_scores_gemma":[0.9940348,0.0027727836,0.00023914724,0.001466257,0.0011422674,0.0003447141],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003475305,0.0016682262,0.0014486925,0.0025232022,0.0007047498,0.0016631522,0.0022797848,0.002168862,0.016045],"category_scores_gemma":[0.016786018,0.0002911779,0.001039562,0.0013819094,0.001091555,0.004447197,0.0048343195,0.0025405528,0.005284428],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093998463,0.0006746372,0.0034628683,0.0005153983,0.00017051352,0.00036667954,0.00052198983,0.03698588,0.025938343,0.025803797,0.042127747,0.86249214],"study_design_scores_gemma":[0.00039457405,0.000806477,0.00311961,0.00011166002,0.00012244614,0.0006699559,0.00028263926,0.8530237,0.021812864,0.082741536,0.03674431,0.0001701753],"about_ca_topic_score_codex":0.0015447366,"about_ca_topic_score_gemma":0.004203944,"teacher_disagreement_score":0.016045,"about_ca_system_score_codex":0.00069169054,"about_ca_system_score_gemma":0.0013650703,"threshold_uncertainty_score":0.05367589},"labels":[],"label_agreement":null},{"id":"W4389524279","doi":"10.18653/v1/2023.newsum-1.8","title":"Generating Extractive and Abstractive Summaries in Parallel from Scientific Articles Incorporating Citing Statements","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Graph; Information retrieval; Citation; Natural language processing; Quality (philosophy); Artificial intelligence; World Wide Web; Theoretical computer science","score_opus":0.06240136843857767,"score_gpt":0.31145803579328124,"score_spread":0.24905666735470355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524279","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049316436,0.00091396336,0.9379998,0.0005455183,0.00025902112,0.0002424536,0.0011674261,0.006946032,0.0026094997],"genre_scores_gemma":[0.31021157,0.001073228,0.6748413,0.00017255003,0.00041932904,0.00034474715,0.0043199924,0.0008672845,0.007749934],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993411,0.00014038062,0.00007409082,0.00016427941,0.000240695,0.00003948783],"domain_scores_gemma":[0.99576205,0.0018836659,0.00041395382,0.0006278152,0.0011655965,0.00014700837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015620415,0.0012605655,0.00096044113,0.0019218306,0.00038894944,0.0017047955,0.0009276837,0.0009110447,0.0022549885],"category_scores_gemma":[0.008725194,0.00044772838,0.0007595613,0.0015374171,0.00026373673,0.0019152099,0.0014150888,0.0012106086,0.0020181735],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048054152,0.00023136304,0.001957419,0.0007417065,0.00024627056,0.00037509398,0.00075433135,0.08277723,0.09309151,0.008287368,0.01102705,0.8000301],"study_design_scores_gemma":[0.00008683968,0.00037663136,0.0020542294,0.000052537736,0.00029761138,0.00018992348,0.00023484825,0.8559103,0.09848998,0.023430435,0.018798305,0.00007839725],"about_ca_topic_score_codex":0.0010281012,"about_ca_topic_score_gemma":0.003122964,"teacher_disagreement_score":0.0022549885,"about_ca_system_score_codex":0.00035890215,"about_ca_system_score_gemma":0.0010210897,"threshold_uncertainty_score":0.008260965},"labels":[],"label_agreement":null},{"id":"W4389524293","doi":"10.18653/v1/2023.nlposs-1.3","title":"Deepparse : An Extendable, and Fine-Tunable State-Of-The-Art Library for Parsing Multinational Street Addresses","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Parsing; Computer science; Python (programming language); Artificial intelligence; Open source; Programming language; Natural language processing; Dependency grammar; Top-down parsing; Segmentation; Software","score_opus":0.03731307456453981,"score_gpt":0.2700880737081976,"score_spread":0.2327749991436578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524293","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0076242844,0.00075270876,0.436389,0.00045521493,0.00025502266,0.0001551252,0.025955282,0.5163624,0.012050958],"genre_scores_gemma":[0.11864114,0.001806339,0.5995484,0.0015655083,0.00018914038,0.0010373527,0.16782694,0.07888843,0.030496744],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99896276,0.00015786714,0.00010016566,0.0003630276,0.00028538675,0.00013082479],"domain_scores_gemma":[0.9988042,0.0003655587,0.00011486432,0.00040068154,0.00023710862,0.00007760073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010611938,0.0024822967,0.00093719136,0.0023183003,0.0007292117,0.0025632652,0.003957448,0.0013840536,0.026052082],"category_scores_gemma":[0.0057449276,0.0015074614,0.0024652376,0.002468071,0.0009578582,0.006642894,0.004493119,0.003697774,0.023376051],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005538366,0.00028109312,0.0057117674,0.0014860365,0.0003866041,0.0004929303,0.0008034946,0.035404205,0.0106720235,0.032951564,0.44794047,0.46331596],"study_design_scores_gemma":[0.00014651314,0.0001413741,0.003235917,0.00036458686,0.00018541337,0.00066951296,0.00025596833,0.4226914,0.048710264,0.078563385,0.4447936,0.00024204682],"about_ca_topic_score_codex":0.0088283,"about_ca_topic_score_gemma":0.016417887,"teacher_disagreement_score":0.026052082,"about_ca_system_score_codex":0.0013259674,"about_ca_system_score_gemma":0.0028441725,"threshold_uncertainty_score":0.0871529},"labels":[],"label_agreement":null},{"id":"W4389524312","doi":"10.18653/v1/2023.nllp-1.26","title":"Tracing Influence at Scale: A Contrastive Learning Approach to Linking Public Comments and Regulator Responses","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ivey Foundation; University of British Columbia","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Tracing; Matching (statistics); Computer science; Set (abstract data type); Scale (ratio); Regulator; Test set; Test (biology); Simple (philosophy); Artificial intelligence; Machine learning; Programming language; Mathematics; Epistemology","score_opus":0.03797563552236556,"score_gpt":0.26544662827626364,"score_spread":0.2274709927538981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524312","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2214394,0.00093792786,0.7635882,0.0014125053,0.00015647769,0.000418388,0.00071547,0.0039859423,0.0073457225],"genre_scores_gemma":[0.86554736,0.00016532395,0.12854409,0.0004358305,0.00020148388,0.00023334722,0.0010579118,0.00020270869,0.0036119781],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99664485,0.0017017117,0.00012431864,0.0009954995,0.00038282544,0.00015086691],"domain_scores_gemma":[0.98153794,0.014575032,0.001078667,0.00124737,0.0013001433,0.0002608075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045223013,0.00115511,0.00079801783,0.0037552032,0.00094790943,0.001687819,0.0022875387,0.0021950253,0.0018082631],"category_scores_gemma":[0.022973882,0.00059084303,0.0010012554,0.0020401166,0.0015555026,0.0032819067,0.0026153082,0.002580283,0.0009487944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011324377,0.0009749691,0.04486513,0.00043322123,0.0005452858,0.00096197514,0.004631494,0.19354221,0.02811121,0.019779941,0.010221345,0.69480073],"study_design_scores_gemma":[0.000037204267,0.0000934724,0.0034359368,0.000018856652,0.00005507955,0.00009004279,0.00019146231,0.97818506,0.0044612032,0.011086824,0.0023166048,0.000028273516],"about_ca_topic_score_codex":0.009610169,"about_ca_topic_score_gemma":0.013828016,"teacher_disagreement_score":0.009610169,"about_ca_system_score_codex":0.0017702634,"about_ca_system_score_gemma":0.0010140887,"threshold_uncertainty_score":0.023916543},"labels":[],"label_agreement":null},{"id":"W4389524336","doi":"10.18653/v1/2023.emnlp-main.942","title":"A State-Vector Framework for Dataset Effects","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Component (thermodynamics); Machine learning; Artificial intelligence; Quality (philosophy); Vector space; Artificial neural network; State (computer science); Data mining; Space (punctuation); Mathematics; Algorithm","score_opus":0.033934969120240195,"score_gpt":0.31167007748815495,"score_spread":0.27773510836791476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524336","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068137557,0.0005134378,0.98690933,0.0013381237,0.000121240555,0.000099515906,0.00068566314,0.0005627206,0.0029562511],"genre_scores_gemma":[0.60233235,0.002889082,0.36961046,0.0013797763,0.0009909668,0.0030556896,0.0035490019,0.0009446074,0.015248158],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9934469,0.0033677442,0.000408641,0.0012938707,0.0009696671,0.0005130791],"domain_scores_gemma":[0.9600417,0.02705517,0.002720831,0.005765891,0.0036818741,0.000734562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016746266,0.0020553363,0.0027920138,0.0027299097,0.0014405374,0.006305494,0.005125985,0.002875816,0.013007531],"category_scores_gemma":[0.04711937,0.0014587062,0.0022213906,0.0028320865,0.004344768,0.012914733,0.0063882633,0.006825409,0.001867381],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020641585,0.00014152631,0.0024371382,0.00022372574,0.0001322062,0.00014251097,0.00020839815,0.24592105,0.0011815435,0.71046466,0.0042082393,0.034732573],"study_design_scores_gemma":[0.000029740455,0.000080417405,0.00037026452,0.000048394937,0.000061003033,0.00003370347,0.00002804784,0.76369447,0.0005560582,0.23231025,0.002738495,0.000049234226],"about_ca_topic_score_codex":0.0071299775,"about_ca_topic_score_gemma":0.005095584,"teacher_disagreement_score":0.016746266,"about_ca_system_score_codex":0.0026359565,"about_ca_system_score_gemma":0.0029343588,"threshold_uncertainty_score":0.0885638},"labels":[],"label_agreement":null},{"id":"W4389524352","doi":"10.18653/v1/2023.findings-emnlp.861","title":"COMET-M: Reasoning about Multiple Events in Complex Sentences","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Comet; Event (particle physics); Computer science; Natural language processing; Coreference; Sentence; Context (archaeology); Inference; Artificial intelligence; Resolution (logic); Meaning (existential); Natural (archaeology); Commonsense knowledge; Natural language; Psychology; Knowledge-based systems; History; Physics","score_opus":0.06975010678736791,"score_gpt":0.2996717027449558,"score_spread":0.22992159595758788,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524352","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03227485,0.0017458741,0.88296807,0.0021402016,0.00044424357,0.00091347535,0.024174692,0.046290606,0.009048043],"genre_scores_gemma":[0.22394793,0.00064128276,0.71638215,0.0014027329,0.00030093463,0.0006100276,0.051523607,0.0014998801,0.0036914805],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997407,0.0005949053,0.00018952238,0.0011882769,0.00051139336,0.00010890621],"domain_scores_gemma":[0.9914471,0.0063094893,0.00040695982,0.0009835124,0.000594011,0.00025893882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025448883,0.0024924218,0.0008163151,0.0027842063,0.0013176649,0.0029393241,0.00452328,0.0028747376,0.010508499],"category_scores_gemma":[0.019496752,0.00096367626,0.0037465044,0.0012290131,0.0010445989,0.0075968048,0.0038614115,0.0041919863,0.00304707],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011296534,0.0007506257,0.015409126,0.0049356506,0.0014362363,0.003406271,0.0035247528,0.088484436,0.025482433,0.06628019,0.1744198,0.6147408],"study_design_scores_gemma":[0.00020684025,0.00017814069,0.003552796,0.00029189058,0.0003559389,0.0013190147,0.0006703142,0.81136507,0.018770266,0.10593274,0.057231683,0.00012531385],"about_ca_topic_score_codex":0.0072428645,"about_ca_topic_score_gemma":0.018960197,"teacher_disagreement_score":0.010508499,"about_ca_system_score_codex":0.0015555226,"about_ca_system_score_gemma":0.0016897733,"threshold_uncertainty_score":0.03515446},"labels":[],"label_agreement":null},{"id":"W4389524397","doi":"10.18653/v1/2023.emnlp-main.100","title":"Increasing Coverage and Precision of Textual Information in Multilingual Knowledge Graphs","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Bridging (networking); Natural language processing; Artificial intelligence; Machine translation; Task (project management); Information retrieval; Knowledge graph; Question answering; Benchmark (surveying); Entity linking; Quality (philosophy); Knowledge base","score_opus":0.020362852201209255,"score_gpt":0.277810869769017,"score_spread":0.2574480175678077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524397","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51928216,0.011239773,0.38150173,0.0037241424,0.0004505695,0.0003376543,0.022624655,0.041228056,0.0196112],"genre_scores_gemma":[0.7644419,0.0020276234,0.17953277,0.00057591434,0.00014881702,0.00016714851,0.04775798,0.002719406,0.0026282952],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9866882,0.0048530856,0.0010161054,0.004563187,0.0024216087,0.00045775133],"domain_scores_gemma":[0.90920925,0.06587023,0.0036109663,0.012644173,0.007817511,0.0008478952],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0091193225,0.001478868,0.0014518566,0.01171833,0.001911518,0.0036945252,0.002312832,0.0021761367,0.0024057133],"category_scores_gemma":[0.0850502,0.00078251504,0.0015671984,0.0077672703,0.002059838,0.011085889,0.005566952,0.0028074658,0.0013623363],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017752948,0.0008302808,0.03703432,0.0067878934,0.0016199476,0.0012040861,0.006295579,0.17342308,0.028854463,0.020260975,0.059891954,0.6620221],"study_design_scores_gemma":[0.0002902247,0.00052865065,0.030534383,0.0009029311,0.001212062,0.0013399757,0.0038781567,0.7052765,0.0983213,0.079624064,0.07778316,0.00030857013],"about_ca_topic_score_codex":0.018190067,"about_ca_topic_score_gemma":0.026639597,"teacher_disagreement_score":0.018190067,"about_ca_system_score_codex":0.0022081966,"about_ca_system_score_gemma":0.0021763896,"threshold_uncertainty_score":0.048228145},"labels":[],"label_agreement":null},{"id":"W4389524398","doi":"10.18653/v1/2023.emnlp-main.63","title":"Tree of Clarifications: Answering Ambiguous Questions with Retrieval-Augmented Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Ministry of Science and ICT, South Korea; National Research Foundation of Korea; Institute for Information and Communications Technology Promotion; Electronics and Telecommunications Research Institute; National Research Foundation","keywords":"Computer science; Question answering; Ambiguity; Tree (set theory); Set (abstract data type); Artificial intelligence; Code (set theory); Natural language processing; Language model; Information retrieval; Domain (mathematical analysis); Machine learning; Mathematics; Programming language","score_opus":0.027444748966846972,"score_gpt":0.2701363584174316,"score_spread":0.24269160945058463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524398","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034827046,0.003792085,0.9118403,0.0018702172,0.00043734486,0.0006543353,0.0037997272,0.03699948,0.005779497],"genre_scores_gemma":[0.3069578,0.00087774976,0.66657317,0.0016199889,0.0003144203,0.00060580386,0.013548193,0.0013249727,0.008177892],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975745,0.0013414085,0.000091389215,0.00061668054,0.00026724572,0.0001088211],"domain_scores_gemma":[0.9941248,0.004045393,0.00022764248,0.00086964655,0.0005102671,0.00022224913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033258356,0.0019642715,0.0011068438,0.0017607395,0.00079523603,0.0020524638,0.0024443492,0.0027245835,0.0064366506],"category_scores_gemma":[0.014551011,0.00067237107,0.0016120502,0.0011605076,0.0008229296,0.005972543,0.0029672976,0.0037558738,0.0046917987],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011382214,0.0006370793,0.0036435658,0.0012935858,0.00033305856,0.00063099945,0.0033925609,0.09828797,0.019628001,0.02570833,0.112979345,0.73232734],"study_design_scores_gemma":[0.00011849492,0.00018646527,0.000751478,0.000101148784,0.000092983784,0.00023969046,0.00061906513,0.9129184,0.00567178,0.052611843,0.026613615,0.00007501336],"about_ca_topic_score_codex":0.008621345,"about_ca_topic_score_gemma":0.017438244,"teacher_disagreement_score":0.008621345,"about_ca_system_score_codex":0.0012855877,"about_ca_system_score_gemma":0.0014500604,"threshold_uncertainty_score":0.021532714},"labels":[],"label_agreement":null},{"id":"W4389524414","doi":"10.18653/v1/2023.emnlp-demo.12","title":"Spacerini: Plug-and-play Search Engines with Pyserini and Hugging Face","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Face (sociological concept); Computer science; Zhàng; Linguistics; History; China; Philosophy; Archaeology","score_opus":0.030535421691699427,"score_gpt":0.26145688013681834,"score_spread":0.2309214584451189,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524414","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010905249,0.0014564238,0.09256046,0.0005476832,0.0006500752,0.000397466,0.01153212,0.86773396,0.014216568],"genre_scores_gemma":[0.31282404,0.0029310123,0.2934947,0.002574586,0.00067091413,0.0023658816,0.11977957,0.16077746,0.10458186],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99927,0.00013087889,0.00004668063,0.00016724346,0.00026529084,0.00011987247],"domain_scores_gemma":[0.99868804,0.00065415486,0.000051373147,0.00027144974,0.00013071026,0.00020415495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001322557,0.0021460648,0.0013519009,0.001230196,0.00080190663,0.0018951884,0.0030904557,0.0011169465,0.042669166],"category_scores_gemma":[0.0048314272,0.0013388672,0.001028887,0.0012047642,0.0005232552,0.005638023,0.0030348024,0.0018583252,0.031406987],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0048883655,0.0005595265,0.0025614912,0.001056185,0.00027827683,0.00046214787,0.00060788146,0.0027249956,0.016863441,0.0102397725,0.76565933,0.1940985],"study_design_scores_gemma":[0.002114537,0.0009248866,0.0042524585,0.00021159097,0.00028684855,0.0008737211,0.0004821843,0.27049723,0.07288132,0.021913566,0.62510645,0.0004552204],"about_ca_topic_score_codex":0.004036053,"about_ca_topic_score_gemma":0.0052878987,"teacher_disagreement_score":0.042669166,"about_ca_system_score_codex":0.00061849644,"about_ca_system_score_gemma":0.0010359525,"threshold_uncertainty_score":0.14274263},"labels":[],"label_agreement":null},{"id":"W4389524452","doi":"10.18653/v1/2023.emnlp-main.12","title":"Sparse Universal Transformer","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Transformer; Computer science; Inference; Computation; Scaling; Language model; Generalization; Theoretical computer science; Artificial intelligence; Algorithm; Computer engineering; Mathematics","score_opus":0.03619836200111303,"score_gpt":0.24093978456339393,"score_spread":0.2047414225622809,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008034252,0.0002003802,0.9827321,0.00025247154,0.000064199834,0.00009618052,0.00043732594,0.003468373,0.004714685],"genre_scores_gemma":[0.49871585,0.00059655594,0.48442286,0.00076927344,0.00015616267,0.00037100716,0.0025277953,0.0012571315,0.0111833075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987758,0.00028390912,0.00008346786,0.00035721646,0.00034604713,0.00015357608],"domain_scores_gemma":[0.99800915,0.00076079083,0.00011916511,0.0007233421,0.00029400244,0.00009351468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014144084,0.0010461052,0.0011649778,0.0008804497,0.0006714589,0.0014769335,0.0023374723,0.0010905224,0.010214819],"category_scores_gemma":[0.0071724528,0.00065138744,0.0018462199,0.0009337549,0.0015975146,0.005629379,0.0034450497,0.0020487616,0.0027369773],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041936664,0.00015436769,0.0020210794,0.00051588344,0.00020321412,0.00044170354,0.0003798557,0.22436862,0.016094932,0.3137631,0.023521395,0.4181166],"study_design_scores_gemma":[0.000033588636,0.00007500176,0.00017676815,0.000028593297,0.00005165646,0.0002525308,0.000049167094,0.774319,0.0065257684,0.20909512,0.009368106,0.000024708248],"about_ca_topic_score_codex":0.0029811265,"about_ca_topic_score_gemma":0.0054208473,"teacher_disagreement_score":0.010214819,"about_ca_system_score_codex":0.0010381998,"about_ca_system_score_gemma":0.002223594,"threshold_uncertainty_score":0.03417194},"labels":[],"label_agreement":null},{"id":"W4389524537","doi":"10.18653/v1/2023.findings-emnlp.822","title":"Efficiently Enhancing Zero-Shot Performance of Instruction Following Model via Retrieval of Soft Prompt","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Embedding; Benchmark (surveying); Generalization; Task (project management); Inference; Zero (linguistics); Shot (pellet); Artificial intelligence; Scaling; Machine learning; Mathematics","score_opus":0.02492057637310018,"score_gpt":0.2507083064115961,"score_spread":0.2257877300384959,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524537","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4137396,0.005152839,0.5119793,0.0016683054,0.0010458803,0.0003848672,0.0026372243,0.05460523,0.0087868],"genre_scores_gemma":[0.9035788,0.00042251727,0.081994124,0.00077575835,0.0001247099,0.00018969656,0.005537067,0.0008099097,0.0065673804],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99903584,0.00024211309,0.000052843112,0.00041563966,0.00010760366,0.00014592441],"domain_scores_gemma":[0.9973605,0.0013729455,0.00009882163,0.00070419733,0.00026800114,0.00019548349],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020054432,0.0023134812,0.001572701,0.0005472077,0.0005122755,0.0016985987,0.0026354766,0.0020933512,0.005536102],"category_scores_gemma":[0.010479915,0.0005432037,0.0011204482,0.0005152895,0.0007840086,0.004887613,0.0025172099,0.0036420324,0.0035389236],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002039659,0.0010956994,0.008909847,0.00087337487,0.00026636047,0.00035023177,0.00054837397,0.19501892,0.028127847,0.00437127,0.026695231,0.7317032],"study_design_scores_gemma":[0.000079085636,0.00046356505,0.0011993173,0.000048768296,0.000057494337,0.00012303898,0.00022798602,0.9743978,0.012417801,0.008624912,0.00231741,0.000042872944],"about_ca_topic_score_codex":0.0055983495,"about_ca_topic_score_gemma":0.007919432,"teacher_disagreement_score":0.0055983495,"about_ca_system_score_codex":0.0009378787,"about_ca_system_score_gemma":0.0015872975,"threshold_uncertainty_score":0.018520117},"labels":[],"label_agreement":null},{"id":"W4389524581","doi":"10.18653/v1/2023.findings-emnlp.1036","title":"Can Retriever-Augmented Language Models Reason? The Blame Game Between the Retriever and the Language Model","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Language model; Natural language processing; Artificial intelligence; Metric (unit)","score_opus":0.023412816619829673,"score_gpt":0.25256404098141105,"score_spread":0.2291512243615814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389524581","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054398045,0.0014466888,0.9041462,0.004202715,0.0002635312,0.00032474825,0.0020264215,0.024021951,0.009169634],"genre_scores_gemma":[0.4135685,0.0008584955,0.5673403,0.0023928676,0.00015994822,0.00033244418,0.0061043557,0.0013329724,0.007910192],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958248,0.0018573899,0.00029116543,0.0010177769,0.0007799931,0.00022895698],"domain_scores_gemma":[0.9878407,0.006805899,0.00046768962,0.0035581274,0.0010803834,0.00024713954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007044787,0.0018162411,0.0013779813,0.0013133417,0.0006472308,0.00382274,0.0041204137,0.0022603301,0.008813233],"category_scores_gemma":[0.028621525,0.0011646354,0.0028395662,0.0008945339,0.0016356154,0.017058888,0.0037030336,0.0045109275,0.005767064],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014286153,0.000664949,0.0060585574,0.0018084829,0.00076277583,0.00073902716,0.002420389,0.1887931,0.02139702,0.078432105,0.0431679,0.6543271],"study_design_scores_gemma":[0.00011258261,0.00020989837,0.00057275395,0.000103094484,0.00017711055,0.00030843104,0.00047242516,0.8804141,0.010425385,0.08851306,0.018595034,0.0000960556],"about_ca_topic_score_codex":0.010970721,"about_ca_topic_score_gemma":0.019627826,"teacher_disagreement_score":0.010970721,"about_ca_system_score_codex":0.0016727446,"about_ca_system_score_gemma":0.0022219995,"threshold_uncertainty_score":0.037256896},"labels":[],"label_agreement":null},{"id":"W4389577465","doi":"10.1109/ichi57859.2023.00103","title":"Leveraging Foundation Models for Clinical Text Analysis","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Public Health Ontario; University of Toronto","funders":"","keywords":"Computer science; Transformer; Information extraction; Data extraction; Data science; Artificial intelligence; Machine learning; Data mining; Natural language processing; Information retrieval; MEDLINE; Engineering","score_opus":0.2348475168951484,"score_gpt":0.4032111303280295,"score_spread":0.16836361343288112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389577465","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018781967,0.00076825847,0.9718911,0.00077372225,0.00008958842,0.00022745198,0.0019809385,0.0035423501,0.00194458],"genre_scores_gemma":[0.53666884,0.0016788926,0.44148055,0.0005050306,0.00031930016,0.0005176899,0.014103005,0.00043967922,0.0042870976],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99892336,0.00034391094,0.00012951507,0.00029630045,0.00023064615,0.00007613218],"domain_scores_gemma":[0.9953951,0.003142034,0.0003229998,0.00040930187,0.00063094904,0.000099595505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023877365,0.00086592947,0.0004719545,0.0028950407,0.00042839485,0.0012858704,0.0010562269,0.00069892243,0.0024449748],"category_scores_gemma":[0.009710636,0.0003368017,0.0012902404,0.001839925,0.00046002562,0.0038732672,0.0013638312,0.0014492936,0.0021962267],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055657834,0.0003279991,0.011954359,0.00059345627,0.00027525154,0.0007078962,0.0006579797,0.0893062,0.01555086,0.028006643,0.021596251,0.83046645],"study_design_scores_gemma":[0.000034571538,0.00009355026,0.0023486055,0.000066327644,0.00010519412,0.00025538477,0.00012381794,0.9376125,0.006614938,0.040079974,0.0126296785,0.000035538436],"about_ca_topic_score_codex":0.0063326866,"about_ca_topic_score_gemma":0.0109036025,"teacher_disagreement_score":0.0063326866,"about_ca_system_score_codex":0.0008310585,"about_ca_system_score_gemma":0.0020982637,"threshold_uncertainty_score":0.012627661},"labels":[],"label_agreement":null},{"id":"W4389613952","doi":"10.1162/coli_a_00501","title":"Stance Detection with Explanations","year":2023,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Automatic summarization; Artificial intelligence; Parsing; Benchmark (surveying); Machine learning; Focus (optics); Grammaticality; Natural language processing; Linguistics","score_opus":0.02863169177005188,"score_gpt":0.26809391849028813,"score_spread":0.23946222672023626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389613952","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20961457,0.0063972846,0.7385594,0.0026578805,0.00046234723,0.0006710536,0.013786606,0.017654693,0.0101961475],"genre_scores_gemma":[0.7088115,0.0008199917,0.27216762,0.00024136434,0.0003297065,0.00019967381,0.01473169,0.00025389914,0.0024444666],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998616,0.00043149356,0.00012332093,0.00032235926,0.00040176362,0.000105036284],"domain_scores_gemma":[0.99119663,0.0052050115,0.001311971,0.00077143934,0.0012927842,0.00022224338],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014555064,0.001054901,0.00055490684,0.0085789375,0.0006047592,0.0013616158,0.0009264894,0.0015860912,0.0027081114],"category_scores_gemma":[0.014499505,0.0003262233,0.0009229287,0.00267564,0.00047285963,0.0024262434,0.0013662721,0.001457186,0.0018410209],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084614847,0.00028005816,0.047922954,0.001342828,0.00031059803,0.0011594258,0.0012107486,0.024347538,0.01736165,0.01881768,0.03806642,0.8483339],"study_design_scores_gemma":[0.00008834033,0.00021426297,0.015914274,0.0004342632,0.00022324873,0.0007338021,0.00067671324,0.8560657,0.024208313,0.06648106,0.03488279,0.0000772199],"about_ca_topic_score_codex":0.0013218989,"about_ca_topic_score_gemma":0.0019064323,"teacher_disagreement_score":0.0085789375,"about_ca_system_score_codex":0.00059509446,"about_ca_system_score_gemma":0.0010078653,"threshold_uncertainty_score":0.009059548},"labels":[],"label_agreement":null},{"id":"W4389630342","doi":"10.1109/icacte59887.2023.10335393","title":"Research on Multi-knowledge Graph and Semantic-aware for Automatic Text Summarization","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina","funders":"Natural Science Foundation of Shanghai","keywords":"Automatic summarization; Computer science; Knowledge graph; Information retrieval; Graph; Natural language processing; Text graph; Artificial intelligence; Entity linking; Embedding; Knowledge base; Theoretical computer science","score_opus":0.15724904843764811,"score_gpt":0.40007169816590155,"score_spread":0.24282264972825343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389630342","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019412315,0.0028068668,0.97214025,0.0006135346,0.00009246249,0.00008077052,0.00035975655,0.0028637575,0.0016303689],"genre_scores_gemma":[0.5352431,0.0046222773,0.44914442,0.0005712839,0.000309542,0.00020767553,0.0031682292,0.0005043429,0.0062291045],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99926704,0.00017565062,0.000055907614,0.0002994459,0.00015217815,0.000049841197],"domain_scores_gemma":[0.9988794,0.00048290458,0.00017137827,0.00017142846,0.00023652207,0.00005844642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007478751,0.001007367,0.0008409863,0.0033293369,0.00047950266,0.0012135694,0.0016963595,0.001089032,0.0013077647],"category_scores_gemma":[0.0026668021,0.00041475255,0.0015110228,0.0032492466,0.0005772997,0.005755696,0.0007621798,0.0010237371,0.0005429209],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018255963,0.00024709245,0.0036577727,0.0009361446,0.0004124024,0.00030724626,0.0005367507,0.25480422,0.022545367,0.046321094,0.008382712,0.6616667],"study_design_scores_gemma":[0.0000142243825,0.00008838848,0.0013333395,0.000032786746,0.00014842884,0.00010195296,0.000093035465,0.9559355,0.005414841,0.028290076,0.008514371,0.000033059438],"about_ca_topic_score_codex":0.009966844,"about_ca_topic_score_gemma":0.011694167,"teacher_disagreement_score":0.009966844,"about_ca_system_score_codex":0.0011675934,"about_ca_system_score_gemma":0.0009929149,"threshold_uncertainty_score":0.01981765},"labels":[],"label_agreement":null},{"id":"W4389923528","doi":"10.21449/ijate.1394194","title":"Language models in automated essay scoring: Insights for the Turkish language","year":2023,"lang":"en","type":"article","venue":"International Journal of Assessment Tools in Education","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Turkish; Transformative learning; Computer science; Language model; Artificial intelligence; Natural language processing; Intersection (aeronautics); Transformer; Linguistics; Sociology; Engineering","score_opus":0.0407797704534206,"score_gpt":0.39349338321123484,"score_spread":0.3527136127578142,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389923528","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56087387,0.001829608,0.40815893,0.0053223073,0.00009407947,0.00015556903,0.00048059385,0.0005325463,0.022552513],"genre_scores_gemma":[0.96676666,0.00022775384,0.03166266,0.00007667572,0.000030466528,0.000045386594,0.00013725461,0.00004370776,0.0010094739],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956204,0.0035116693,0.00013729776,0.00027492552,0.0003478659,0.000107799126],"domain_scores_gemma":[0.9829196,0.013808279,0.0010234237,0.0006548136,0.0013622552,0.00023173571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041127484,0.00045092576,0.00034520368,0.0014497007,0.0005865339,0.003842706,0.00064413546,0.0005667455,0.001729318],"category_scores_gemma":[0.024966124,0.00021494308,0.00041694657,0.0012795365,0.0010416211,0.004149931,0.0013650202,0.0011479935,0.0006754607],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005091795,0.00030945975,0.052394904,0.00046196915,0.00014042233,0.0011471998,0.016004875,0.16363727,0.006191288,0.33971298,0.0055290046,0.41396144],"study_design_scores_gemma":[0.000022137809,0.0001222591,0.010400089,0.0001166214,0.000046999434,0.00041938323,0.004952389,0.7994181,0.0016920537,0.16885717,0.013877441,0.00007542494],"about_ca_topic_score_codex":0.007036024,"about_ca_topic_score_gemma":0.009395199,"teacher_disagreement_score":0.007036024,"about_ca_system_score_codex":0.0016897974,"about_ca_system_score_gemma":0.0017548074,"threshold_uncertainty_score":0.02175057},"labels":[],"label_agreement":null},{"id":"W4389945749","doi":"10.1109/wi-iat59888.2023.00044","title":"PPPG-DialoGPT: A Prompt-based and Personality-aware Framework For Conversational Recommendation Systems","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Personality; Recommender system; Human–computer interaction; Multimedia; World Wide Web; Psychology; Social psychology","score_opus":0.06840691760978909,"score_gpt":0.3033723702233297,"score_spread":0.2349654526135406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389945749","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012897998,0.00060063833,0.9603682,0.0002704752,0.00008528178,0.0004958729,0.001340137,0.022103433,0.0018378673],"genre_scores_gemma":[0.25454357,0.0003721433,0.73435557,0.00026716545,0.000081863676,0.00073366764,0.003441827,0.0004781676,0.0057261055],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99874866,0.0004184939,0.00011480362,0.00042009278,0.00020645288,0.000091534916],"domain_scores_gemma":[0.9987728,0.0004343637,0.00008389544,0.0002803508,0.00025108494,0.0001774233],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015253708,0.0012333861,0.00084728515,0.00075416703,0.0006774008,0.001211074,0.002002793,0.0015368909,0.0032500343],"category_scores_gemma":[0.004682002,0.0005396154,0.0012757314,0.0004788527,0.00047157495,0.002132628,0.0018188329,0.0019942874,0.0020081524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017429086,0.0012240745,0.011992707,0.0015368959,0.00034381394,0.0014274482,0.0026967756,0.14363073,0.053419944,0.030109648,0.041916564,0.70995855],"study_design_scores_gemma":[0.00007712698,0.0003024009,0.0018245982,0.00004346721,0.00007595152,0.00033179607,0.00018807559,0.9501632,0.008311119,0.014377387,0.024219118,0.000085721775],"about_ca_topic_score_codex":0.01256412,"about_ca_topic_score_gemma":0.01744187,"teacher_disagreement_score":0.01256412,"about_ca_system_score_codex":0.0008885784,"about_ca_system_score_gemma":0.0017889565,"threshold_uncertainty_score":0.024981976},"labels":[],"label_agreement":null},{"id":"W4389945798","doi":"10.1109/wi-iat59888.2023.00023","title":"Triple Extraction with Generative Technique for Constructing Weighted Knowledge Graph","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Knowledge graph; Computer science; Transformer; Relationship extraction; Encoder; Graph; Generative grammar; Artificial intelligence; Natural language processing; Theoretical computer science; Relation (database); Data mining","score_opus":0.03649364172802901,"score_gpt":0.29641802285979507,"score_spread":0.2599243811317661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389945798","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045620203,0.00015092158,0.98975676,0.00010557549,0.000029391298,0.00014468326,0.0011017943,0.0029987688,0.0011501323],"genre_scores_gemma":[0.13911343,0.00057719275,0.8436457,0.00019564269,0.000045080673,0.00039752587,0.011608516,0.0007352312,0.0036816392],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999074,0.00016113756,0.00008167994,0.00036326543,0.0002643693,0.000055492776],"domain_scores_gemma":[0.9985607,0.0006500859,0.00010747466,0.0003599799,0.0002839766,0.000037737263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007593207,0.0011079969,0.00063700875,0.0042646835,0.00082372746,0.0008898464,0.0012735418,0.00077037205,0.0040951404],"category_scores_gemma":[0.0042617605,0.0006117527,0.0018622484,0.0032192129,0.0006192314,0.0025974074,0.0019010992,0.0014313816,0.0026098262],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018420568,0.00015947208,0.0029545124,0.0007246374,0.00020641965,0.0015080045,0.0010349449,0.044703715,0.033927258,0.06269272,0.02041613,0.8314879],"study_design_scores_gemma":[0.000045198467,0.00010464076,0.0015072934,0.00014317039,0.00027977827,0.0010830495,0.00039465222,0.6999536,0.058083136,0.1978229,0.04049375,0.000088839144],"about_ca_topic_score_codex":0.0049587004,"about_ca_topic_score_gemma":0.009531655,"teacher_disagreement_score":0.0049587004,"about_ca_system_score_codex":0.00067694904,"about_ca_system_score_gemma":0.0016202836,"threshold_uncertainty_score":0.013699591},"labels":[],"label_agreement":null},{"id":"W4389966905","doi":"10.1017/9781108974004.017","title":"Great Expectations and EPIC Fails: A Computational Perspective on Irony and Sarcasm","year":2023,"lang":"en","type":"book-chapter","venue":"Cambridge University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Irony; Sarcasm; Utterance; Perspective (graphical); Computer science; Meaning (existential); EPIC; Linguistics; Psychology; Artificial intelligence; Epistemology; Literature; Philosophy; Art","score_opus":0.035023779391520266,"score_gpt":0.22428555701681094,"score_spread":0.18926177762529067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389966905","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11910034,0.003414787,0.35110614,0.01792499,0.00034165816,0.0000722912,0.00025327702,0.00064520136,0.5071414],"genre_scores_gemma":[0.9168532,0.0008002484,0.05496894,0.0008575151,0.0001273393,0.00009366864,0.00023425279,0.00022090711,0.02584387],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986859,0.0008276297,0.00003645914,0.00015397492,0.0002101485,0.000085968364],"domain_scores_gemma":[0.99700207,0.0024689427,0.00012676309,0.00018826859,0.00012760267,0.00008628496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017647285,0.0005661376,0.00041607206,0.0010935819,0.0018792679,0.004484626,0.0014784822,0.0018237534,0.006325438],"category_scores_gemma":[0.005856631,0.0005288102,0.00078844355,0.00073392456,0.007878138,0.007442869,0.0017623293,0.0025854807,0.00075742137],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010488916,0.000010612787,0.00034040931,0.000027024826,0.000006208405,0.00009004126,0.0026647171,0.0026297772,0.00012359169,0.98741734,0.0015746352,0.005105169],"study_design_scores_gemma":[0.0000072424014,0.000013883247,0.0004149902,0.000056620236,0.000009254201,0.00016168614,0.0015263071,0.026011148,0.00022519086,0.95573145,0.015828043,0.0000141354985],"about_ca_topic_score_codex":0.0028415935,"about_ca_topic_score_gemma":0.0032792015,"teacher_disagreement_score":0.006325438,"about_ca_system_score_codex":0.0020778775,"about_ca_system_score_gemma":0.00086298556,"threshold_uncertainty_score":0.021160662},"labels":[],"label_agreement":null},{"id":"W4389986588","doi":"10.52098/acj.2023346","title":"&lt;b&gt;Systematic Review of Semantic Analysis Methods&lt;/b&gt;","year":2023,"lang":"en","type":"article","venue":"Applied computing Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Computer science; Latent semantic analysis; Semantics (computer science); Semantic analysis (machine learning); Natural language processing; Artificial intelligence; Subcategory; Data science; Programming language","score_opus":0.030778487012252904,"score_gpt":0.31685856909619564,"score_spread":0.28608008208394275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389986588","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004848579,0.99200714,0.0022093153,0.0030075584,0.00085593085,0.0003094432,0.0003316243,0.00002043564,0.00077367696],"genre_scores_gemma":[0.010885958,0.974603,0.0087864585,0.0030120006,0.00064598024,0.0012000849,0.00040398553,0.000029804462,0.00043275917],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9692645,0.015341929,0.0076352926,0.001992884,0.0053650313,0.00040046894],"domain_scores_gemma":[0.85776085,0.10392916,0.014917401,0.00429809,0.018224895,0.0008695697],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.032828864,0.0016375624,0.0035207574,0.021849122,0.001166943,0.0036153947,0.001896672,0.0020257498,0.0052846866],"category_scores_gemma":[0.1297174,0.0010390328,0.0044831056,0.017905787,0.0021704033,0.0064032143,0.0024185358,0.0019609064,0.0009760323],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015516314,0.000052447005,0.0012503394,0.6617953,0.0036547922,0.00012934998,0.00096157205,0.00031882955,0.00042862186,0.005999387,0.02279371,0.30246046],"study_design_scores_gemma":[0.000106739644,0.00021732233,0.0041132066,0.8134277,0.008692794,0.00047252933,0.0012807804,0.00033102764,0.0005411848,0.00621091,0.16452995,0.000075883145],"about_ca_topic_score_codex":0.0055779093,"about_ca_topic_score_gemma":0.017430035,"teacher_disagreement_score":0.96717113,"about_ca_system_score_codex":0.004436852,"about_ca_system_score_gemma":0.02485402,"threshold_uncertainty_score":0.17361772},"labels":[],"label_agreement":null},{"id":"W4390094923","doi":"10.1145/3570945.3607324","title":"Intent and Entity Detection with Data Augmentation for a Mental Health Virtual Assistant Chatbot","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Chatbot; Computer science; Conversation; Resource (disambiguation); Interdependence; Data science; Mental health; World Wide Web; Information retrieval; Medicine","score_opus":0.07853758234781592,"score_gpt":0.31863874441462986,"score_spread":0.24010116206681392,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390094923","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36084932,0.0007328244,0.5710989,0.0020068095,0.00060351676,0.0016424338,0.0047104806,0.053791225,0.0045643877],"genre_scores_gemma":[0.46625084,0.00011010554,0.51794076,0.00040193106,0.00009686349,0.0011988481,0.009828145,0.00044263995,0.0037299623],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954181,0.0021785533,0.00033904085,0.001021851,0.0007819367,0.00026055102],"domain_scores_gemma":[0.98651296,0.009472374,0.0004385741,0.0013182124,0.001699379,0.000558532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060476195,0.0015563256,0.00094902,0.0026404746,0.0012361241,0.0015574382,0.0024428512,0.0017180012,0.002548271],"category_scores_gemma":[0.01381831,0.00069613644,0.00095059094,0.000970538,0.0006848681,0.0027592382,0.0026299637,0.0021663888,0.0019770316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002569319,0.0034727212,0.024845915,0.0017249665,0.00027979125,0.0019750958,0.0038083845,0.037434183,0.12127531,0.0034682234,0.029966502,0.7691796],"study_design_scores_gemma":[0.0001338135,0.00090440345,0.011954339,0.00009034893,0.00009879858,0.0005726812,0.0009398664,0.9041704,0.05813281,0.00381393,0.019042464,0.0001462095],"about_ca_topic_score_codex":0.0047299424,"about_ca_topic_score_gemma":0.007196707,"teacher_disagreement_score":0.0060476195,"about_ca_system_score_codex":0.0008779492,"about_ca_system_score_gemma":0.0012097067,"threshold_uncertainty_score":0.031983256},"labels":[],"label_agreement":null},{"id":"W4390097263","doi":"10.1109/jbhi.2023.3346210","title":"EHR-HGCN: An Enhanced Hybrid Approach for Text Classification Using Heterogeneous Graph Convolutional Networks in Electronic Health Records","year":2023,"lang":"en","type":"article","venue":"IEEE Journal of Biomedical and Health Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"National Natural Science Foundation of China; Natural Science Foundation of Jilin Province","keywords":"Computer science; Sentence; Artificial intelligence; Natural language processing; Graph; Convolutional neural network; Graph database; Text graph; Biomedical text mining; Text mining; Information retrieval; Theoretical computer science","score_opus":0.08133374493038933,"score_gpt":0.34122620279012134,"score_spread":0.25989245785973203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390097263","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14773864,0.001808121,0.8286915,0.0014751375,0.0002894622,0.00029655217,0.0027151525,0.0123519115,0.004633531],"genre_scores_gemma":[0.59667295,0.00082943856,0.374269,0.0009847496,0.00025049812,0.00025713694,0.010863387,0.0004730388,0.015399752],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994511,0.00011563985,0.000035271154,0.00018882289,0.00012701315,0.00008219314],"domain_scores_gemma":[0.99931884,0.00023976248,0.000078846315,0.00013019062,0.00018956068,0.000042842166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073544594,0.0010385276,0.0005292804,0.0019999954,0.00047547414,0.00063348,0.0014142145,0.000986683,0.0017305491],"category_scores_gemma":[0.0021754217,0.0002759705,0.000802446,0.0016040945,0.0003734868,0.001966027,0.0010182745,0.0011286094,0.00077431847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043114307,0.0004674509,0.0070626987,0.00022755226,0.000265588,0.00037676166,0.0002597606,0.1712984,0.01533686,0.008224493,0.023533508,0.77251583],"study_design_scores_gemma":[0.000015369978,0.000050006995,0.0011629697,0.00001343545,0.00004344302,0.00004662209,0.00003721538,0.98514014,0.0042115333,0.0061412356,0.0031252466,0.000012871072],"about_ca_topic_score_codex":0.019373814,"about_ca_topic_score_gemma":0.031649582,"teacher_disagreement_score":0.019373814,"about_ca_system_score_codex":0.0011073147,"about_ca_system_score_gemma":0.0012556926,"threshold_uncertainty_score":0.038522065},"labels":[],"label_agreement":null},{"id":"W4390136468","doi":"10.48550/arxiv.2312.13454","title":"MixEHR-SurG: a joint proportional hazard and guided topic model for inferring mortality-associated topics from electronic health records","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Health records; Hazard; Joint (building); Health hazard; Electronic health record; Computer science; Data science; Data mining; Medicine; Environmental health; Engineering; Political science; Health care; Civil engineering; Biology","score_opus":0.2372352507708456,"score_gpt":0.25520962357286203,"score_spread":0.01797437280201644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390136468","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.083966844,0.0020518454,0.9034182,0.0014469452,0.00018292444,0.00040838876,0.0045761215,0.0028193877,0.0011293517],"genre_scores_gemma":[0.6803985,0.0016699908,0.29084316,0.0010072999,0.0010258333,0.0012259639,0.01611391,0.0004078083,0.0073075774],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978807,0.0011677486,0.000112758076,0.00056141784,0.0001715706,0.00010586622],"domain_scores_gemma":[0.9918108,0.006944464,0.00038401663,0.00035977,0.00036056657,0.00014028451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006886895,0.0013224977,0.0013706478,0.0018169563,0.00057164824,0.0011050812,0.002094827,0.0013632352,0.002286826],"category_scores_gemma":[0.010339056,0.0006511947,0.002615695,0.0012496433,0.0006344229,0.0016838129,0.0016034495,0.0023094148,0.0012700423],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001621096,0.00073890184,0.061692625,0.00071694556,0.0012726156,0.0005816452,0.0017582412,0.443493,0.004743378,0.025953598,0.022869222,0.4345588],"study_design_scores_gemma":[0.00009148294,0.00008189396,0.0026562468,0.000031573203,0.00010122743,0.000089302484,0.00006851958,0.97964066,0.00051962363,0.014416258,0.0022695307,0.000033606677],"about_ca_topic_score_codex":0.0084366705,"about_ca_topic_score_gemma":0.012388841,"teacher_disagreement_score":0.0084366705,"about_ca_system_score_codex":0.00084244524,"about_ca_system_score_gemma":0.0013407309,"threshold_uncertainty_score":0.036421835},"labels":[],"label_agreement":null},{"id":"W4390188528","doi":"10.1109/dasc/picom/cbdcom/cy59711.2023.10361304","title":"Unveiling Uncertainty: Supporting Learners Through NLP-Driven Confusion Identification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Computer science; Artificial intelligence; Confusion; Popularity; World Wide Web; Identification (biology); Natural language processing; Deep learning; Task (project management); Machine learning; Data science; Engineering","score_opus":0.04950785522572413,"score_gpt":0.3185036625247137,"score_spread":0.26899580729898953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390188528","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22446202,0.00090798916,0.75475043,0.0027070232,0.00017269536,0.0005229955,0.0013828494,0.009385278,0.005708684],"genre_scores_gemma":[0.7997158,0.0002477882,0.1953246,0.0004365444,0.00013301578,0.0003274587,0.0018812803,0.0003003078,0.0016332747],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99574995,0.0020177842,0.00026553427,0.0009439683,0.0007675473,0.00025525372],"domain_scores_gemma":[0.9761979,0.018608313,0.001337669,0.0014619516,0.0017729965,0.0006211748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006832379,0.0019428993,0.0009922386,0.002520977,0.0013257249,0.003945261,0.0020277568,0.0019296363,0.0020638057],"category_scores_gemma":[0.042884998,0.00052756513,0.001054489,0.0010057986,0.0010916882,0.008129321,0.006351022,0.0039436063,0.0014537154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013048578,0.0013750307,0.06901908,0.00086553127,0.00024362381,0.00070479227,0.012267209,0.07815369,0.018335491,0.010431439,0.010750279,0.7965491],"study_design_scores_gemma":[0.000042516418,0.00019103027,0.0036672773,0.000102654274,0.000089322406,0.00019734286,0.002647731,0.93189746,0.014188944,0.03955317,0.007342586,0.00007997849],"about_ca_topic_score_codex":0.0038914904,"about_ca_topic_score_gemma":0.0048712255,"teacher_disagreement_score":0.006832379,"about_ca_system_score_codex":0.0011254477,"about_ca_system_score_gemma":0.0022553436,"threshold_uncertainty_score":0.03613347},"labels":[],"label_agreement":null},{"id":"W4390215202","doi":"10.48550/arxiv.2312.14769","title":"Large Language Model (LLM) Bias Index -- LLMBI","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Canada West","funders":"","keywords":"Metric (unit); Computer science; Index (typography); Measure (data warehouse); Empirical measure; Reliability (semiconductor); Econometrics; Data science; Gender bias; Performance metric; Artificial intelligence; Cognitive psychology; Machine learning; Natural language processing; Psychology; Data mining; Statistics; Social psychology; Mathematics; Economics; Operations management","score_opus":0.15653195457383753,"score_gpt":0.2146151699651807,"score_spread":0.058083215391343174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390215202","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12815915,0.0029170483,0.8216779,0.0038686702,0.00078488083,0.001489278,0.014473152,0.0072403084,0.019389534],"genre_scores_gemma":[0.70391124,0.0007463274,0.27425852,0.0014779874,0.00037571992,0.0021171195,0.012459496,0.0015185793,0.0031349584],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97017384,0.015052341,0.0030146465,0.0031134237,0.007884259,0.0007615739],"domain_scores_gemma":[0.88802093,0.08033445,0.007781233,0.011024186,0.011834073,0.0010051079],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03188123,0.0016611962,0.0015210711,0.0039459164,0.0014225107,0.004782772,0.0019803888,0.0019236762,0.0047890837],"category_scores_gemma":[0.14857598,0.00045542096,0.0016688069,0.003973681,0.0014675565,0.005999019,0.004597037,0.0032457614,0.0024859195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019272342,0.0004414953,0.13726686,0.0034067945,0.0022607564,0.0006810941,0.004946992,0.13722534,0.014889434,0.07433155,0.061553612,0.56106883],"study_design_scores_gemma":[0.00016144317,0.0008512241,0.034622602,0.0008818707,0.0006521472,0.0010059385,0.0022879941,0.69431156,0.02021715,0.17996316,0.06453805,0.0005066954],"about_ca_topic_score_codex":0.0028271044,"about_ca_topic_score_gemma":0.003625871,"teacher_disagreement_score":0.03188123,"about_ca_system_score_codex":0.0023331463,"about_ca_system_score_gemma":0.0027199646,"threshold_uncertainty_score":0.16860604},"labels":[],"label_agreement":null},{"id":"W4390231356","doi":"10.18280/ria.370622","title":"Combined Approach for Answer Identification with Small Sized Reading Comprehension Datasets","year":2023,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Identification (biology); Reading comprehension; Reading (process); Computer science; Comprehension; Natural language processing; Artificial intelligence; Information retrieval; Psychology; Linguistics; Programming language; Philosophy; Biology","score_opus":0.09111594763010084,"score_gpt":0.2880590670980732,"score_spread":0.19694311946797238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390231356","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.131983,0.001026468,0.8314856,0.00067947793,0.00015003509,0.00096851785,0.007052857,0.02250748,0.004146608],"genre_scores_gemma":[0.40453744,0.00020406421,0.5699618,0.00020121035,0.0001380886,0.001184108,0.020734137,0.00039070533,0.002648491],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99724776,0.0010472955,0.00026168747,0.0008575052,0.00044862556,0.0001370699],"domain_scores_gemma":[0.99451786,0.0028264124,0.00021560866,0.0010312684,0.0012691835,0.00013964354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031247535,0.0014426962,0.0012722058,0.0048980047,0.0006659593,0.001710949,0.0017851995,0.0018264584,0.0043833726],"category_scores_gemma":[0.011271726,0.0002924577,0.001386219,0.0023271125,0.00031301338,0.0034515532,0.002482429,0.0014598095,0.004080366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007148701,0.0010041348,0.012728245,0.00062673946,0.00042213616,0.0004675864,0.0007488445,0.014047687,0.039632127,0.0027409417,0.012989144,0.9138774],"study_design_scores_gemma":[0.00013221895,0.00060070475,0.014610564,0.000056220335,0.00027575146,0.0005612319,0.0014749957,0.9162787,0.033343237,0.014142145,0.018422425,0.0001017584],"about_ca_topic_score_codex":0.002563742,"about_ca_topic_score_gemma":0.0049180957,"teacher_disagreement_score":0.0048980047,"about_ca_system_score_codex":0.00060600037,"about_ca_system_score_gemma":0.0010689242,"threshold_uncertainty_score":0.016525507},"labels":[],"label_agreement":null},{"id":"W4390327836","doi":"10.1109/imcet59736.2023.10368218","title":"Cost-Effective Tweet Classification through Transfer Learning in Low-Resource NLP Settings","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Agence Universitaire de la Francophonie","keywords":"Computer science; Transformer; Natural language processing; Artificial intelligence; Transfer of learning; Machine learning; Labeled data","score_opus":0.047659383262126404,"score_gpt":0.2929655279400818,"score_spread":0.2453061446779554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390327836","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3152976,0.002201296,0.6407011,0.0021927701,0.00035793486,0.00053680595,0.003410323,0.02598705,0.009315106],"genre_scores_gemma":[0.85433555,0.0004029601,0.1302417,0.0004059153,0.00016597426,0.00037690118,0.0057797604,0.0003048562,0.007986406],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988003,0.0004844535,0.00007132645,0.0003561216,0.00014765152,0.00014018401],"domain_scores_gemma":[0.9967213,0.0021952882,0.0001481934,0.00047699924,0.00036353758,0.00009467469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026644266,0.0017105679,0.0010645052,0.0015319567,0.00086631195,0.0013126582,0.0021235198,0.0017729595,0.0032724768],"category_scores_gemma":[0.0074069635,0.00042251733,0.0009040245,0.0015244576,0.0006540134,0.003985191,0.0019946224,0.0019680695,0.0032653857],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088161905,0.0007861552,0.005900518,0.00039184562,0.00021778005,0.00040542235,0.00031823164,0.23823696,0.010349146,0.0043170513,0.017279595,0.7209156],"study_design_scores_gemma":[0.000034195313,0.00009005598,0.0006754257,0.000010221095,0.000023641696,0.000042123458,0.00011886585,0.98691016,0.0044994014,0.0062495177,0.0013304626,0.000015950505],"about_ca_topic_score_codex":0.009067321,"about_ca_topic_score_gemma":0.009012964,"teacher_disagreement_score":0.009067321,"about_ca_system_score_codex":0.0013540179,"about_ca_system_score_gemma":0.0013976836,"threshold_uncertainty_score":0.018029094},"labels":[],"label_agreement":null},{"id":"W4390338861","doi":"10.1101/2023.12.25.23300520","title":"Empowering Transformers for Evidence-Based Medicine","year":2023,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs","keywords":"Transformer; Computer science; Clinical Practice; Bootstrapping (finance); Artificial intelligence; Medicine; Engineering; Family medicine","score_opus":0.15222435695892006,"score_gpt":0.3598493581918209,"score_spread":0.20762500123290084,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390338861","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0075102937,0.0022022685,0.9807011,0.0033405675,0.00013640207,0.00019068635,0.00057181664,0.0021592837,0.003187637],"genre_scores_gemma":[0.38154644,0.002637031,0.6099974,0.001089263,0.00047451374,0.00031931011,0.0015177069,0.000300222,0.0021181058],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9905033,0.0057612206,0.0007646766,0.0008163792,0.001988959,0.00016542371],"domain_scores_gemma":[0.9310272,0.05533536,0.0025055935,0.0061058765,0.0041627283,0.00086323975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013083297,0.0008409354,0.0008865061,0.004933784,0.0004190259,0.0030843432,0.0012764217,0.0010952705,0.007534312],"category_scores_gemma":[0.073394045,0.00052120636,0.0013626452,0.002874747,0.0015139652,0.005373417,0.0038554058,0.0025227328,0.0021857605],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005081111,0.00018418928,0.004254946,0.0022939325,0.00036089827,0.00027240394,0.0007238197,0.061894804,0.0071387542,0.118917026,0.012966036,0.7904851],"study_design_scores_gemma":[0.00006857415,0.00016900487,0.0011532272,0.0006647056,0.0001341869,0.00024418347,0.00021258197,0.4648795,0.009751283,0.49065265,0.032020316,0.00004981758],"about_ca_topic_score_codex":0.0012061726,"about_ca_topic_score_gemma":0.001311591,"teacher_disagreement_score":0.013083297,"about_ca_system_score_codex":0.0011830892,"about_ca_system_score_gemma":0.0024432915,"threshold_uncertainty_score":0.06919187},"labels":[],"label_agreement":null},{"id":"W4390405403","doi":"10.18280/ts.400644","title":"Automated Generation of Chinese Text-Image Summaries Using Deep Learning Techniques","year":2023,"lang":"en","type":"article","venue":"Traitement du signal","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Deep learning; Image (mathematics); Natural language processing; Pattern recognition (psychology); Information retrieval","score_opus":0.03644585053335536,"score_gpt":0.2878999087316401,"score_spread":0.25145405819828476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390405403","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043393347,0.0005591349,0.9485946,0.00020814687,0.000106354724,0.00011141565,0.00038073055,0.0051571066,0.0014891532],"genre_scores_gemma":[0.48289073,0.00056822406,0.50642365,0.00012318451,0.00013393018,0.0001795383,0.0022413146,0.00044566015,0.006993772],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9997608,0.000038762137,0.000023773631,0.00008183921,0.00006752498,0.000027302833],"domain_scores_gemma":[0.9993917,0.00019652133,0.00008701332,0.000087943685,0.00019801506,0.00003888353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00049960957,0.00090905995,0.0005631555,0.0009827614,0.00030400956,0.00060609047,0.0007003149,0.0004693108,0.0019988995],"category_scores_gemma":[0.0015938265,0.00023123909,0.00067282264,0.0007348063,0.00021688809,0.0011289353,0.0005462554,0.00065330294,0.0008805647],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041624022,0.00012424035,0.001278101,0.00036549265,0.00008817223,0.00035420895,0.00033463613,0.118783556,0.074720755,0.007862461,0.009040036,0.7866321],"study_design_scores_gemma":[0.000024196814,0.00011233612,0.00055949565,0.000007875891,0.000043952332,0.000069453316,0.000056759225,0.9604186,0.031675804,0.0035441802,0.0034716493,0.000015765121],"about_ca_topic_score_codex":0.0026014636,"about_ca_topic_score_gemma":0.0045402157,"teacher_disagreement_score":0.0026014636,"about_ca_system_score_codex":0.0004968956,"about_ca_system_score_gemma":0.00059215265,"threshold_uncertainty_score":0.0066869855},"labels":[],"label_agreement":null},{"id":"W4390437770","doi":"10.48550/arxiv.2312.16917","title":"Unified Lattice Graph Fusion for Chinese Named Entity Recognition","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Natural language processing; Artificial intelligence; Lexicon; Named-entity recognition; Graph; Adjacency list; Leverage (statistics); Adjacency matrix; Theoretical computer science; Task (project management); Algorithm","score_opus":0.13955799615271985,"score_gpt":0.2165424273676652,"score_spread":0.07698443121494536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390437770","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026578782,0.00063353626,0.96489054,0.00031050533,0.00007312871,0.000068058165,0.00057227723,0.0043306635,0.0025425584],"genre_scores_gemma":[0.6303381,0.000773032,0.35616758,0.00042177003,0.00011762435,0.00022934891,0.005657239,0.00051883253,0.005776503],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914384,0.00022116714,0.00003983585,0.00029507192,0.0001988807,0.00010119494],"domain_scores_gemma":[0.9992799,0.00026458653,0.000057586865,0.00017115667,0.00018176321,0.000045040677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009654495,0.0009807864,0.0011112706,0.0021191048,0.00069115544,0.0009309793,0.0015060833,0.0009498414,0.0026757368],"category_scores_gemma":[0.002464541,0.00036618928,0.0013454349,0.0029089774,0.00068514165,0.0033364568,0.0019102187,0.0013604925,0.0013618356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024724187,0.00020101332,0.0020549868,0.00019231724,0.00014856832,0.00019993616,0.00026991524,0.2893544,0.01582919,0.03489286,0.01504824,0.64156127],"study_design_scores_gemma":[0.000009547428,0.000022498689,0.00031781374,0.0000053450235,0.000019287547,0.000027990462,0.000035541434,0.9712371,0.0030859518,0.023217166,0.0020077685,0.000013924133],"about_ca_topic_score_codex":0.014141673,"about_ca_topic_score_gemma":0.018940933,"teacher_disagreement_score":0.014141673,"about_ca_system_score_codex":0.0011582187,"about_ca_system_score_gemma":0.0015018422,"threshold_uncertainty_score":0.02811873},"labels":[],"label_agreement":null},{"id":"W4390542458","doi":"10.1007/s13369-023-08567-1","title":"WASM: A Dataset for Hashtag Recommendation for Arabic Tweets","year":2024,"lang":"en","type":"article","venue":"Arabian Journal for Science and Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Social media; Information retrieval; Arabic; Task (project management); Categorization; Benchmark (surveying); Microblogging; Natural language processing; Artificial intelligence; World Wide Web","score_opus":0.0375174136116534,"score_gpt":0.2986276063840521,"score_spread":0.2611101927723987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390542458","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023430092,0.00088862574,0.0015128782,0.00037269742,0.0002984562,0.00028386698,0.9662493,0.0033211587,0.0036430005],"genre_scores_gemma":[0.016849818,0.00028186745,0.0051022717,0.00012382733,0.000073617644,0.00031331915,0.974096,0.00010342885,0.0030558573],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993303,0.00011873304,0.0000929595,0.00014536036,0.00019677836,0.00011594023],"domain_scores_gemma":[0.998747,0.00029042704,0.00012100567,0.00020994301,0.00044515138,0.00018654384],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004866338,0.0018561073,0.00080495566,0.00516275,0.0010733037,0.0010194125,0.0009365621,0.0016498028,0.010091682],"category_scores_gemma":[0.0032632821,0.0002924672,0.0008355738,0.004240854,0.0002574847,0.0011319473,0.0010994085,0.0009606846,0.018876618],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082765444,0.00043067575,0.01440071,0.0024722747,0.0002081613,0.00050868234,0.00056265923,0.0023890794,0.00936232,0.0012823695,0.9037172,0.06383822],"study_design_scores_gemma":[0.00037505952,0.00035950233,0.06866156,0.0005587113,0.00021376586,0.000894155,0.002125083,0.024264423,0.012183538,0.0020593011,0.88804036,0.0002645793],"about_ca_topic_score_codex":0.023909602,"about_ca_topic_score_gemma":0.04118613,"teacher_disagreement_score":0.023909602,"about_ca_system_score_codex":0.0009820696,"about_ca_system_score_gemma":0.0015557115,"threshold_uncertainty_score":0.047540843},"labels":[],"label_agreement":null},{"id":"W4390573228","doi":"10.48550/arxiv.2401.00907","title":"LaFFi: Leveraging Hybrid Natural Language Feedback for Fine-tuning Language Models","year":2023,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; University of Alberta; Alberta Machine Intelligence Institute; Compute Canada; Mitacs; Canadian Institute for Advanced Research","keywords":"Computer science; Fine-tuning; Realm; Natural language; Natural (archaeology); Language model; Natural language understanding; Task (project management); Artificial intelligence; Natural language processing; Engineering","score_opus":0.1021838348557336,"score_gpt":0.21279216872488352,"score_spread":0.11060833386914992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390573228","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06347111,0.0025747458,0.8297341,0.0014690417,0.0005334514,0.00040211627,0.002405089,0.09575179,0.0036586854],"genre_scores_gemma":[0.50308007,0.00055293454,0.47195283,0.0021188448,0.00041566687,0.0010104256,0.010378991,0.0032913284,0.0071988283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99671125,0.0014706032,0.00014633764,0.0010946218,0.000399018,0.00017808875],"domain_scores_gemma":[0.9892902,0.0071963905,0.0003775581,0.0018164273,0.0010251252,0.0002944733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056453366,0.0028486212,0.0013682734,0.0014665673,0.000644299,0.0015628262,0.0035069862,0.0024475965,0.003233025],"category_scores_gemma":[0.024481885,0.0007126066,0.0013767626,0.00082907156,0.000779486,0.0043382626,0.0025721681,0.0039302274,0.0035387806],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089001004,0.00094776007,0.005667846,0.0006992581,0.00045924928,0.00034141235,0.0008557862,0.18765078,0.04312557,0.003971949,0.0426764,0.712714],"study_design_scores_gemma":[0.00009235513,0.00015160335,0.0004976855,0.00002572102,0.000040340463,0.00006946782,0.000049305694,0.9787756,0.008721614,0.007321467,0.0042075394,0.000047262558],"about_ca_topic_score_codex":0.006543498,"about_ca_topic_score_gemma":0.014599704,"teacher_disagreement_score":0.006543498,"about_ca_system_score_codex":0.0013009555,"about_ca_system_score_gemma":0.0015043969,"threshold_uncertainty_score":0.029855728},"labels":[],"label_agreement":null},{"id":"W4390632530","doi":"10.48550/arxiv.2401.02297","title":"Are LLMs Robust for Spoken Dialogues?","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; McGill University","keywords":"Perplexity; Robustness (evolution); Spoken language; Computer science; Task (project management); Natural language processing; Artificial intelligence; Set (abstract data type); Speech recognition; Language model","score_opus":0.17638339622166002,"score_gpt":0.1969008740361831,"score_spread":0.020517477814523066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390632530","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16029765,0.0058023296,0.76814234,0.0034600238,0.0012335789,0.00033322533,0.0041280314,0.047192723,0.009410089],"genre_scores_gemma":[0.8514921,0.00091233914,0.12817068,0.0014348566,0.00033348895,0.0004718402,0.006939932,0.003973594,0.0062712645],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99139905,0.004189638,0.00040343936,0.0023836363,0.0010845219,0.0005396301],"domain_scores_gemma":[0.9814528,0.012349996,0.0008008355,0.0033309646,0.0016731715,0.00039212417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069880015,0.0022147472,0.0017342633,0.0011506923,0.00061319047,0.0031587307,0.0020280425,0.0022447277,0.0045129233],"category_scores_gemma":[0.04851823,0.0007266783,0.0011395246,0.0006809482,0.0011733884,0.00333922,0.0027010196,0.0029496762,0.007401003],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001551639,0.0002783434,0.008856121,0.0015621737,0.0008687541,0.00048068984,0.0015996778,0.22732763,0.056744102,0.004009412,0.02410986,0.6726116],"study_design_scores_gemma":[0.00014648923,0.0004166405,0.007596218,0.00026791677,0.00020169873,0.00046781066,0.0007799392,0.90845764,0.049833123,0.018369166,0.01329452,0.00016881594],"about_ca_topic_score_codex":0.0052566067,"about_ca_topic_score_gemma":0.004988668,"teacher_disagreement_score":0.0069880015,"about_ca_system_score_codex":0.00082261796,"about_ca_system_score_gemma":0.0013611646,"threshold_uncertainty_score":0.03695655},"labels":[],"label_agreement":null},{"id":"W4390789641","doi":"10.1017/s1351324923000542","title":"Lightweight transformers for clinical natural language processing","year":2024,"lang":"en","type":"article","venue":"Natural Language Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"NIHR Imperial Biomedical Research Centre; Instituto de Salud Carlos III; Medical Research Council; National Institutes of Health; Kementerian Kesihatan Malaysia; All-India Institute of Medical Sciences; Horizon 2020 Framework Programme; Prince Charles Hospital Foundation; Foreign, Commonwealth and Development Office; National Institute for Health Research Health Protection Research Unit; University of Oxford; Norges Forskningsråd; Imperial College London; Conselho Nacional de Desenvolvimento Científico e Tecnológico; University of Cape Town; Public Health England; Wellcome Trust; University College Dublin; Sunnybrook Research Institute; Institut National de la Santé et de la Recherche Médicale; Canadian Institutes of Health Research; National Institute for Health and Care Research; Ministero della Salute; European Federation of Pharmaceutical Industries and Associations; European Commission; Bill and Melinda Gates Foundation","keywords":"Computer science; Transformer; Natural language processing; Artificial intelligence; Electrical engineering","score_opus":0.011690292682715328,"score_gpt":0.30424274043960325,"score_spread":0.29255244775688793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390789641","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01791121,0.0013522463,0.8986432,0.0024453534,0.00032328002,0.0005248201,0.0059841727,0.06904554,0.003770186],"genre_scores_gemma":[0.37173125,0.0012735588,0.59954345,0.0014416703,0.0002141247,0.0006950929,0.016596658,0.002284962,0.0062192585],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99769956,0.0007538392,0.00026675235,0.0005486269,0.0006125403,0.000118684984],"domain_scores_gemma":[0.99324965,0.0037991249,0.00034989297,0.0016103536,0.00077499886,0.00021602071],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025475745,0.0013141092,0.000578846,0.0016842915,0.00034939847,0.00183358,0.002084572,0.00090729375,0.013515056],"category_scores_gemma":[0.015458226,0.00071737455,0.0013196762,0.0012548575,0.0010818698,0.0061303945,0.0037677381,0.0031775904,0.008818846],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010626093,0.0004112517,0.0043865633,0.0010206095,0.00017790166,0.0005135506,0.00048023282,0.1242555,0.01213633,0.055408567,0.06832554,0.73182124],"study_design_scores_gemma":[0.00012118371,0.0002119466,0.00047860344,0.0001226599,0.00006660144,0.00036045216,0.00014898332,0.8164868,0.015935369,0.12368374,0.042323884,0.000059856327],"about_ca_topic_score_codex":0.0031337882,"about_ca_topic_score_gemma":0.00571743,"teacher_disagreement_score":0.013515056,"about_ca_system_score_codex":0.0015541998,"about_ca_system_score_gemma":0.0029153815,"threshold_uncertainty_score":0.045212388},"labels":[],"label_agreement":null},{"id":"W4390878767","doi":"10.3899/jrheum.2023-0998","title":"Role of Creation of Plain Language Summaries to Disseminate COVID-19 Research Findings to Patients With Rheumatic Diseases","year":2024,"lang":"en","type":"letter","venue":"The Journal of Rheumatology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Canadian Arthritis Patient Alliance","funders":"Pfizer; Teva Pharmaceutical Industries; AstraZeneca; Eli Lilly and Company; Amgen","keywords":"Dissemination; Medicine; Coronavirus disease 2019 (COVID-19); Pandemic; Information Dissemination; Health literacy; Plain language; Public health; Disease; MEDLINE; 2019-20 coronavirus outbreak; Rheumatic disease; Scientific literacy; Scientific evidence; Literacy; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Family medicine; Alternative medicine; Infectious disease (medical specialty); Health care; Pathology; Psychology; World Wide Web; Linguistics; Outbreak","score_opus":0.01754471033385703,"score_gpt":0.31032601717100494,"score_spread":0.2927813068371479,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390878767","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011135152,0.0007381523,0.0009846672,0.9864791,0.0051114103,0.00005565315,0.00010067704,0.00018740246,0.0052293465],"genre_scores_gemma":[0.07240827,0.004504283,0.012554455,0.8257078,0.06764582,0.0007859584,0.00035644096,0.0004566939,0.0155802565],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97449356,0.018149659,0.0031683184,0.0007214416,0.002548354,0.0009186847],"domain_scores_gemma":[0.6662166,0.26004055,0.014200431,0.012919016,0.030119149,0.016504202],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.033339128,0.0004123008,0.0008824684,0.0016364324,0.0031085396,0.0057268976,0.0016811226,0.013131007,0.014505352],"category_scores_gemma":[0.26082048,0.00066461036,0.00087119883,0.00088888506,0.0017642019,0.0063856454,0.0037805822,0.014081372,0.012731197],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012020295,0.00007182995,0.0031770254,0.00016285872,0.000016810958,0.0009552123,0.0016410802,0.00019373567,0.00024685674,0.0035830133,0.90341544,0.08641597],"study_design_scores_gemma":[0.00014549273,0.0002019013,0.0030326135,0.0013320163,0.000050240003,0.002251455,0.0019768989,0.0027949577,0.0006431352,0.0076608183,0.9798038,0.00010664817],"about_ca_topic_score_codex":0.0031862843,"about_ca_topic_score_gemma":0.00458756,"teacher_disagreement_score":0.9942731,"about_ca_system_score_codex":0.0044936906,"about_ca_system_score_gemma":0.007361845,"threshold_uncertainty_score":0.17631626},"labels":[],"label_agreement":null},{"id":"W4390962525","doi":"10.48550/arxiv.2401.07927","title":"Are self-explanations from Large Language Models faithful?","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Counterfactual thinking; Counterfactual conditional; Interpretability; Consistency (knowledge bases); Measure (data warehouse); Set (abstract data type); Task (project management); Computer science; Inference; Cognitive psychology; Linguistics; Psychology; Artificial intelligence; Natural language processing; Social psychology; Data mining; Programming language; Economics","score_opus":0.06590286736689008,"score_gpt":0.19258703477358027,"score_spread":0.1266841674066902,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390962525","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17414805,0.00095528946,0.811556,0.0046349033,0.00014918925,0.00015640269,0.0010313013,0.00310579,0.0042631505],"genre_scores_gemma":[0.93166,0.00025457644,0.0639343,0.0008973294,0.0001587705,0.00012096723,0.0014412446,0.0004971054,0.0010356432],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913663,0.004909725,0.00047267202,0.0015724647,0.0013077406,0.00037091383],"domain_scores_gemma":[0.88299704,0.08448327,0.008193073,0.019603977,0.0035787662,0.0011438375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012340168,0.0010004515,0.00097429584,0.0018363091,0.00080636865,0.0045985826,0.0020499805,0.0018109849,0.0033652578],"category_scores_gemma":[0.11071406,0.0010587435,0.0016400669,0.00082830916,0.0026416555,0.008982058,0.00305693,0.0043331883,0.00088030845],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021789803,0.00050152425,0.100178644,0.001450624,0.0023015651,0.0012043207,0.008052746,0.21879053,0.012317781,0.30197868,0.014824917,0.33621964],"study_design_scores_gemma":[0.00008546882,0.00008437445,0.004669781,0.00013900812,0.00014037211,0.00019985407,0.000352805,0.51563555,0.0037173813,0.47172272,0.0031770864,0.000075572854],"about_ca_topic_score_codex":0.0036918796,"about_ca_topic_score_gemma":0.0045569246,"teacher_disagreement_score":0.012340168,"about_ca_system_score_codex":0.0016599962,"about_ca_system_score_gemma":0.0017239995,"threshold_uncertainty_score":0.06526184},"labels":[],"label_agreement":null},{"id":"W4390966511","doi":"10.2139/ssrn.4678265","title":"Beware of Botshit: How to Manage the Epistemic Risks of Generative Chatbots","year":2024,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; University of Alberta","funders":"","keywords":"Chatbot; Ignorance; Computer science; Generative grammar; Typology; Epistemology; Artificial intelligence; Sociology; Philosophy","score_opus":0.032841139861076014,"score_gpt":0.2910288670453174,"score_spread":0.2581877271842414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390966511","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04140135,0.00029863138,0.92603064,0.0058361166,0.00023330854,0.00040142803,0.00008376111,0.011522328,0.014192398],"genre_scores_gemma":[0.63439184,0.00023344612,0.35243383,0.00073097437,0.00018848767,0.00041978044,0.00019098252,0.002075136,0.00933553],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9897049,0.00606571,0.00045913554,0.0011218712,0.0017092347,0.00093915797],"domain_scores_gemma":[0.93168247,0.03854713,0.00275329,0.017390028,0.0056144567,0.004012657],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.016348239,0.0013621288,0.0014893181,0.0018416703,0.0032904686,0.008266886,0.0038898934,0.0056038448,0.011079112],"category_scores_gemma":[0.082901,0.0013579077,0.0013277846,0.0009270965,0.0034451592,0.020512043,0.009701442,0.0068349647,0.0036349765],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012119632,0.0011793491,0.018092344,0.0007650918,0.00032185912,0.0010169583,0.019214898,0.05792601,0.01669057,0.37269565,0.03205114,0.47883412],"study_design_scores_gemma":[0.00015642533,0.00026186422,0.0014771926,0.00025001762,0.00021863784,0.000634099,0.0036859056,0.47856003,0.0075797723,0.4716522,0.035340868,0.00018297063],"about_ca_topic_score_codex":0.003251437,"about_ca_topic_score_gemma":0.0040660035,"teacher_disagreement_score":0.9967095,"about_ca_system_score_codex":0.0013567273,"about_ca_system_score_gemma":0.0038259935,"threshold_uncertainty_score":0.0864588},"labels":[],"label_agreement":null},{"id":"W4391013074","doi":"10.48550/arxiv.2401.08898","title":"Bridging State and History Representations: Understanding Self-Predictive RL","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; McGill University; DeepMind; Nvidia","keywords":"Bridging (networking); State (computer science); Computer science; Algorithm; Computer security","score_opus":0.12094346956607983,"score_gpt":0.20132425055422498,"score_spread":0.08038078098814515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391013074","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018730486,0.00022578858,0.9761159,0.0013406605,0.000027342428,0.000022596616,0.000072692215,0.0002033942,0.0032611203],"genre_scores_gemma":[0.7905242,0.0005107281,0.20434973,0.00045796437,0.00011229699,0.0001893098,0.00021764605,0.00016753358,0.003470486],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993051,0.00033981437,0.000034514116,0.00013769945,0.0001244138,0.000058427628],"domain_scores_gemma":[0.99472356,0.0040118615,0.00040152838,0.0004426226,0.00024259223,0.00017786957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020692232,0.00048440788,0.00078400265,0.0007117596,0.00045273107,0.0018509342,0.0017247307,0.0015094897,0.0024370141],"category_scores_gemma":[0.012653899,0.000532565,0.00069380115,0.0006522118,0.0021247598,0.0050328365,0.0021766175,0.0023544333,0.00030664596],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003384683,0.00003337338,0.0011656985,0.00008542665,0.0000410131,0.00009023007,0.0005094891,0.54577315,0.00062713283,0.41962966,0.0014338098,0.03057712],"study_design_scores_gemma":[0.0000035027967,0.000004918718,0.000060957205,0.000010132336,0.0000034876605,0.000007848464,0.000021384936,0.8489602,0.0001104104,0.15032321,0.0004896247,0.0000043803257],"about_ca_topic_score_codex":0.0030801613,"about_ca_topic_score_gemma":0.0030228687,"teacher_disagreement_score":0.0030801613,"about_ca_system_score_codex":0.0013851383,"about_ca_system_score_gemma":0.0009820163,"threshold_uncertainty_score":0.010943234},"labels":[],"label_agreement":null},{"id":"W4391050404","doi":"10.2196/52482","title":"Efficient Machine Reading Comprehension for Health Care Applications: Algorithm Development and Validation of a Context Extraction Approach","year":2024,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Context (archaeology); Benchmark (surveying); Correctness; Inference; Process (computing); Artificial intelligence; Sentence; Machine learning; Task (project management); Domain (mathematical analysis); Natural language processing; Algorithm; Mathematics","score_opus":0.08304430576899183,"score_gpt":0.42509251176969104,"score_spread":0.3420482060006992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391050404","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09260131,0.0012619661,0.88952005,0.00077811937,0.00010839458,0.00069069857,0.00087641797,0.0125586055,0.0016044872],"genre_scores_gemma":[0.26261517,0.0003426427,0.7322679,0.00025811975,0.00006797229,0.00093422114,0.0022047611,0.00025795744,0.0010512846],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986659,0.0004976039,0.00012954272,0.00046332955,0.00016063289,0.00008306998],"domain_scores_gemma":[0.991947,0.0062971003,0.00024741492,0.00038154962,0.0010008186,0.00012618232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030880107,0.0018452703,0.0009861981,0.0017521239,0.0005601894,0.0011749163,0.0020374728,0.002343857,0.0031190815],"category_scores_gemma":[0.013663176,0.0005430963,0.001235652,0.0010921448,0.00043686669,0.0018938702,0.0013487603,0.0022478548,0.0015760267],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005525655,0.0005971461,0.011248483,0.00048697827,0.00023392786,0.00036026927,0.00029998628,0.28834242,0.009779612,0.0023570166,0.00577261,0.679969],"study_design_scores_gemma":[0.000040831335,0.00005957156,0.000789522,0.000017285942,0.000021679982,0.000040219507,0.000039509807,0.994232,0.0029510274,0.0012120277,0.0005885819,0.000007748118],"about_ca_topic_score_codex":0.009012951,"about_ca_topic_score_gemma":0.008480824,"teacher_disagreement_score":0.009012951,"about_ca_system_score_codex":0.0015763209,"about_ca_system_score_gemma":0.0019741603,"threshold_uncertainty_score":0.01792097},"labels":[],"label_agreement":null},{"id":"W4391096444","doi":"10.1109/bigdata59044.2023.10386854","title":"Integrating a PICO Clinical Questioning to the QL4POMR Framework for Building Evidence-Based Clinical Case Reports","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Harmonization; Flexibility (engineering); Clinical Practice; MEDLINE; Data science; Medical education; Medicine; Family medicine","score_opus":0.25199225263696445,"score_gpt":0.47810432010929765,"score_spread":0.2261120674723332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391096444","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038117208,0.00044944687,0.9397921,0.0037621155,0.00026079078,0.002571869,0.010227027,0.03155826,0.007566615],"genre_scores_gemma":[0.026589883,0.00030738756,0.95654666,0.000763547,0.000113735354,0.0012427631,0.010207304,0.0018754625,0.0023532095],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9854182,0.0066616912,0.0032013464,0.0020510813,0.002304975,0.00036263172],"domain_scores_gemma":[0.95064086,0.031693034,0.0033145878,0.007628181,0.005240705,0.0014827147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02246121,0.0016234345,0.00087298366,0.009862085,0.001083049,0.00724734,0.0026631828,0.0018777791,0.01759124],"category_scores_gemma":[0.069505386,0.0012080278,0.002907303,0.0034676243,0.0019038355,0.00700229,0.008056354,0.0022892437,0.0075935805],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065617566,0.0004141326,0.009061854,0.008670775,0.00033153742,0.0024978835,0.012246087,0.020339554,0.013453518,0.15540239,0.10470209,0.6722239],"study_design_scores_gemma":[0.00023245113,0.00025795732,0.0028746277,0.003048241,0.00025622424,0.0019497618,0.0030979954,0.09507276,0.016693152,0.195394,0.6808383,0.0002845918],"about_ca_topic_score_codex":0.0041838847,"about_ca_topic_score_gemma":0.0070618126,"teacher_disagreement_score":0.02246121,"about_ca_system_score_codex":0.002458012,"about_ca_system_score_gemma":0.004150796,"threshold_uncertainty_score":0.11878765},"labels":[],"label_agreement":null},{"id":"W4391124783","doi":"10.48550/arxiv.2401.10825","title":"Recent Advances in Named Entity Recognition: A Comprehensive Survey and Comparative Study","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Collège Boréal","funders":"","keywords":"Substring; Computer science; Implementation; Variety (cybernetics); Artificial intelligence; Transformer; Data science; Named-entity recognition; Graph; Natural language processing; Machine learning; Data structure; Theoretical computer science; Task (project management); Software engineering","score_opus":0.2442692213471044,"score_gpt":0.2589056125410715,"score_spread":0.014636391193967063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391124783","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012770383,0.79256535,0.16811517,0.0046404176,0.0018636265,0.00017477547,0.0029088883,0.0028862099,0.014075201],"genre_scores_gemma":[0.06930509,0.76349753,0.13629478,0.0025032193,0.0056188316,0.00019785001,0.016974056,0.0009899074,0.004618682],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9943539,0.0014711308,0.00062588736,0.0018418427,0.0014582631,0.00024899404],"domain_scores_gemma":[0.9695285,0.021850074,0.0010339643,0.0031448237,0.0039421353,0.0005005846],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009783936,0.002387508,0.00332248,0.013130054,0.0010676642,0.003921216,0.0035444195,0.0017498594,0.0045984713],"category_scores_gemma":[0.025832454,0.0010960653,0.0020879393,0.02127384,0.0013660347,0.01651581,0.0033852293,0.0028690966,0.005792514],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015670311,0.00014299857,0.00576654,0.006712139,0.00029647656,0.00011987512,0.0003493591,0.0057989173,0.0011498984,0.010299478,0.039517518,0.92969],"study_design_scores_gemma":[0.000038025522,0.00033066855,0.01611681,0.003918412,0.00090502447,0.0019509371,0.001569563,0.060373142,0.007468532,0.038450148,0.8685772,0.000301595],"about_ca_topic_score_codex":0.0030208533,"about_ca_topic_score_gemma":0.0031900955,"teacher_disagreement_score":0.013130054,"about_ca_system_score_codex":0.0012376102,"about_ca_system_score_gemma":0.0022395593,"threshold_uncertainty_score":0.05174303},"labels":[],"label_agreement":null},{"id":"W4391157657","doi":"10.48550/arxiv.2401.11323","title":"Identifying and Analyzing Performance-Critical Tokens in Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Natural Science Foundation of China; Canadian Institute for Advanced Research","keywords":"Computer science; Task (project management); Encoding (memory); Security token; Context (archaeology); Natural language processing; Artificial intelligence","score_opus":0.08114905561307885,"score_gpt":0.22541048307872458,"score_spread":0.14426142746564574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391157657","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49738216,0.00070041814,0.4948432,0.00070669333,0.00009304256,0.000098201795,0.00058994227,0.0034437703,0.0021425162],"genre_scores_gemma":[0.9470614,0.0001254834,0.050923496,0.00008128937,0.000030015088,0.00008070732,0.00078021124,0.00025490185,0.0006624542],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99782205,0.0010854556,0.00009749946,0.0004910528,0.0003073701,0.00019650949],"domain_scores_gemma":[0.977375,0.017644726,0.0014887482,0.0021233037,0.00076190796,0.0006062885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003705596,0.0009712569,0.0008427817,0.00081395975,0.00044737788,0.0013328869,0.0011854746,0.0010557204,0.0017168557],"category_scores_gemma":[0.033586737,0.00045687062,0.0005251161,0.0006430148,0.0009965027,0.0034767832,0.0024157786,0.0024161797,0.0006655593],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015232179,0.00062787574,0.05530625,0.0013507903,0.00026510408,0.00090483634,0.0021205787,0.3858501,0.09728026,0.029533656,0.007576752,0.4176606],"study_design_scores_gemma":[0.000023204871,0.00024299872,0.005131822,0.000034021825,0.00003099977,0.00011157593,0.0002286385,0.94571626,0.016589193,0.03040502,0.0014497275,0.00003651788],"about_ca_topic_score_codex":0.0017920032,"about_ca_topic_score_gemma":0.0030713445,"teacher_disagreement_score":0.003705596,"about_ca_system_score_codex":0.0007924205,"about_ca_system_score_gemma":0.0010432081,"threshold_uncertainty_score":0.019597292},"labels":[],"label_agreement":null},{"id":"W4391157705","doi":"10.48550/arxiv.2401.11373","title":"Finding a Needle in the Adversarial Haystack: A Targeted Paraphrasing Approach For Uncovering Edge Cases with Minimal Distribution Distortion","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Adversarial system; Generalizability theory; Computer science; Artificial intelligence; Exploit; Classifier (UML); Language model; Machine learning; Generator (circuit theory); Natural language processing; Mathematics","score_opus":0.07982276981754907,"score_gpt":0.20244046572022145,"score_spread":0.12261769590267238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391157705","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023738226,0.00023770801,0.97230667,0.000531205,0.000026493106,0.00007785042,0.00007395084,0.0010041668,0.0020036665],"genre_scores_gemma":[0.7366451,0.00031200048,0.25628442,0.0008974947,0.00011577751,0.00025789186,0.00041185025,0.000467311,0.0046082153],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980558,0.0008695857,0.000092064525,0.00036301845,0.00045988185,0.0001596541],"domain_scores_gemma":[0.9932841,0.0044986974,0.00056069164,0.0011162882,0.00032859258,0.00021164629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029624274,0.0011986278,0.0012883653,0.0011958093,0.0007822191,0.0013814587,0.001846474,0.0020978446,0.0026538873],"category_scores_gemma":[0.014820624,0.0007144633,0.0012011229,0.0007282874,0.0023964278,0.0031779276,0.0033914126,0.0029389956,0.0012646671],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035636924,0.0002636623,0.003381108,0.00025084612,0.00015310668,0.0010343818,0.00058429525,0.69417405,0.018283045,0.091217555,0.00841181,0.18188986],"study_design_scores_gemma":[0.000012323293,0.000048905607,0.0000880783,0.000015285743,0.0000089670375,0.0001153037,0.00002421234,0.9629082,0.0024597032,0.033504453,0.00080472353,0.000009793891],"about_ca_topic_score_codex":0.0009706504,"about_ca_topic_score_gemma":0.0013419853,"teacher_disagreement_score":0.0029624274,"about_ca_system_score_codex":0.00091324566,"about_ca_system_score_gemma":0.0010063133,"threshold_uncertainty_score":0.015666962},"labels":[],"label_agreement":null},{"id":"W4391343109","doi":"10.1109/cicn59264.2023.10402210","title":"The Future of Document Retrieval: Harnessing the Power of OpenAI and AWS Kendra","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Seneca Polytechnic","funders":"","keywords":"Relevance (law); Computer science; Artificial intelligence; World Wide Web; Information retrieval; Transformative learning; Adaptability","score_opus":0.016327241839581785,"score_gpt":0.267426194183712,"score_spread":0.2510989523441302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391343109","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059530526,0.03600361,0.73831266,0.03630313,0.0013519326,0.00022790866,0.0005138701,0.011960508,0.11579582],"genre_scores_gemma":[0.391991,0.014618676,0.5439562,0.0036327224,0.0021738426,0.00018863684,0.0012800811,0.0022454718,0.039913304],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9952527,0.0014801689,0.0003813403,0.0005684577,0.0020643745,0.0002528957],"domain_scores_gemma":[0.9812352,0.008923836,0.00062327954,0.00492822,0.003143648,0.0011458664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008223029,0.00053314195,0.0009509351,0.0055014105,0.0019372734,0.015372694,0.002441227,0.0013980523,0.006567708],"category_scores_gemma":[0.029461058,0.0006525649,0.0006684904,0.005204418,0.0045582573,0.03595865,0.0077417963,0.0030527404,0.0038484747],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003575645,0.00011041162,0.0021400352,0.0003974949,0.00006364651,0.0002022304,0.0031255202,0.001958086,0.0075151394,0.23417254,0.019469876,0.73048735],"study_design_scores_gemma":[0.00009576133,0.00017814194,0.0019871723,0.00033582817,0.00010501098,0.00079269125,0.0018647915,0.0738599,0.008762047,0.52817553,0.3836021,0.00024109548],"about_ca_topic_score_codex":0.005887462,"about_ca_topic_score_gemma":0.005514133,"teacher_disagreement_score":0.015372694,"about_ca_system_score_codex":0.0016034356,"about_ca_system_score_gemma":0.002167033,"threshold_uncertainty_score":0.043488026},"labels":[],"label_agreement":null},{"id":"W4391384425","doi":"10.1136/bjsports-2023-concussion.8","title":"2.7 Does the SCAT5 10-word list improve the distribution of scores over the SCAT3 5-word list in professional hockey players?","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Rehabilitation Institute; University of Toronto","funders":"","keywords":"Word (group theory); Ceiling effect; Word list; Computer science; Mathematics; Medicine; Artificial intelligence; Pathology","score_opus":0.012331541049213338,"score_gpt":0.2673508503134723,"score_spread":0.2550193092642589,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391384425","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99488056,0.00034234844,0.0007639952,0.00025654226,0.00008780408,0.00012961094,0.00044021127,0.000072910945,0.0030260317],"genre_scores_gemma":[0.9937835,0.0002381467,0.0017508359,0.00026840338,0.000055631423,0.00023366841,0.00076226715,0.000023528164,0.00288397],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9987821,0.0002667086,0.00015860377,0.00020259204,0.00042996174,0.0001600498],"domain_scores_gemma":[0.9972066,0.0006317658,0.0008581897,0.00025063552,0.00071499404,0.0003377757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030724618,0.00051641156,0.0005451451,0.0005112657,0.00039622208,0.00097274914,0.0005791238,0.00073333766,0.0074702324],"category_scores_gemma":[0.008070309,0.00030231383,0.000806964,0.00032639533,0.00064838346,0.0008486551,0.0006025712,0.0004309817,0.0024986393],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.024625363,0.0037838393,0.66635,0.0010494284,0.0005346972,0.00021876085,0.0010705756,0.0009857104,0.010922459,0.00043679148,0.0055443966,0.28447807],"study_design_scores_gemma":[0.00085259334,0.024369825,0.9547315,0.000389749,0.00049414113,0.00047835187,0.00093814835,0.0015084133,0.010879718,0.0007416479,0.004543567,0.00007242352],"about_ca_topic_score_codex":0.0046831956,"about_ca_topic_score_gemma":0.011195268,"teacher_disagreement_score":0.0074702324,"about_ca_system_score_codex":0.00052723184,"about_ca_system_score_gemma":0.0010723957,"threshold_uncertainty_score":0.02499038},"labels":[],"label_agreement":null},{"id":"W4391480596","doi":"10.21203/rs.3.rs-3882757/v1","title":"mCodeGPT: Enhancing Cancer Research through Zero-Shot Information Extraction from Clinical Free Text Data","year":2024,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Zero (linguistics); Shot (pellet); Extraction (chemistry); Information extraction; Computer science; Cancer; Text messaging; Information retrieval; Medicine; World Wide Web; Linguistics; Materials science; Chromatography; Chemistry; Philosophy; Internal medicine","score_opus":0.4635973813089228,"score_gpt":0.5652224047004077,"score_spread":0.10162502339148494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391480596","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.045088988,0.013755522,0.81814015,0.004000551,0.0012581624,0.0012989523,0.058088806,0.051913377,0.006455441],"genre_scores_gemma":[0.1316284,0.0034679559,0.77365625,0.00084916316,0.0009247556,0.0010559096,0.08009613,0.002085413,0.006236076],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970319,0.0010587906,0.0002590591,0.000860964,0.00062590843,0.00016331144],"domain_scores_gemma":[0.9877999,0.008855525,0.00042184177,0.0013228115,0.0012272918,0.0003727277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043036644,0.001722956,0.0012347496,0.010703544,0.0009008255,0.002542846,0.0016650566,0.0021944197,0.004899578],"category_scores_gemma":[0.018923175,0.0007088971,0.0019093606,0.006184312,0.00055409106,0.0031633584,0.0029877964,0.0018942226,0.0045390567],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088204787,0.00043948056,0.0076496666,0.0030870377,0.00082489534,0.00058463804,0.0007898837,0.007617787,0.029506674,0.005549534,0.12094518,0.8221231],"study_design_scores_gemma":[0.0006492828,0.0009943031,0.020437906,0.00091693585,0.0020743054,0.0024874595,0.0010486736,0.5184645,0.09354809,0.084276006,0.27477786,0.00032470317],"about_ca_topic_score_codex":0.0049726698,"about_ca_topic_score_gemma":0.009102267,"teacher_disagreement_score":0.010703544,"about_ca_system_score_codex":0.00068911747,"about_ca_system_score_gemma":0.0029011832,"threshold_uncertainty_score":0.022760212},"labels":[],"label_agreement":null},{"id":"W4391520039","doi":"10.1186/s40537-023-00842-0","title":"Survey of transformers and towards ensemble learning using transformers for natural language processing","year":2024,"lang":"en","type":"article","venue":"Journal Of Big Data","topic":"Topic Modeling","field":"Computer Science","cited_by":112,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Carleton University","keywords":"Computer science; Transformer; Automatic summarization; Artificial intelligence; Language model; Question answering; Natural language processing; Classifier (UML); Natural language; Machine learning; Sentiment analysis; Natural language understanding; Ensemble forecasting","score_opus":0.14319068897057724,"score_gpt":0.35099681994531884,"score_spread":0.2078061309747416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391520039","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011972409,0.025399938,0.9537063,0.0012633934,0.00029971782,0.00010042384,0.00029666224,0.0022790593,0.004682138],"genre_scores_gemma":[0.5181679,0.07933551,0.3833855,0.0017671611,0.0013417915,0.00046601315,0.0036895997,0.0007853341,0.011061164],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99861073,0.00046816992,0.00011388969,0.000364155,0.0003644769,0.00007852579],"domain_scores_gemma":[0.99714535,0.0016530094,0.00010031082,0.00045887838,0.0005463647,0.00009612323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028671818,0.0016224958,0.00129751,0.0023993156,0.0004586726,0.0018718132,0.0020855037,0.00087849103,0.002567156],"category_scores_gemma":[0.0066964403,0.00067364395,0.0017425225,0.0031415455,0.00072564714,0.00590084,0.0019836607,0.0023347295,0.0014128287],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019183078,0.00021065756,0.0042485613,0.0007025637,0.0003189522,0.00010122721,0.00022056668,0.10667882,0.0026685824,0.056333028,0.01414452,0.81418073],"study_design_scores_gemma":[0.000015040808,0.00021639015,0.0008394908,0.00011344636,0.000116112096,0.00017610054,0.0000831174,0.90625215,0.0027611686,0.071310304,0.01807844,0.0000382828],"about_ca_topic_score_codex":0.005278216,"about_ca_topic_score_gemma":0.0038880797,"teacher_disagreement_score":0.005278216,"about_ca_system_score_codex":0.0011853713,"about_ca_system_score_gemma":0.0015017656,"threshold_uncertainty_score":0.015163243},"labels":[],"label_agreement":null},{"id":"W4391556990","doi":"10.4108/eai.18-12-2023.2348180","title":"AI-driven Generation of News Summaries Leveraging GPT and Pegasus Summarizer for Efficient Information Extraction","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"International Development Research Centre; Styrelsen för Internationellt Utvecklingssamarbete","keywords":"Computer science; Information extraction; Automatic summarization; Extraction (chemistry); Information retrieval; Natural language generation; Artificial intelligence; Natural language processing; Natural language","score_opus":0.037771114246217825,"score_gpt":0.27545003907752913,"score_spread":0.2376789248313113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391556990","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041256618,0.00065933215,0.93064183,0.0005862085,0.00025099664,0.00038791125,0.0020028884,0.019356407,0.0048577823],"genre_scores_gemma":[0.35832152,0.00046951757,0.6273132,0.00023390085,0.00026750326,0.00048624852,0.005833352,0.0008090481,0.0062658153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995117,0.00012879889,0.000047073037,0.00014125582,0.00013517575,0.00003604578],"domain_scores_gemma":[0.9970606,0.0017018542,0.00018779583,0.0003073981,0.00065562496,0.00008669288],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009621915,0.0009965408,0.00072379457,0.0020766074,0.0004630357,0.0012897673,0.0008832903,0.00074258854,0.0038779727],"category_scores_gemma":[0.0062950486,0.0003323241,0.0006016696,0.0015033352,0.0002211048,0.0015519467,0.00069412787,0.0010219695,0.0023101699],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004776658,0.00033088456,0.0029605706,0.0006434841,0.00019528503,0.00062283326,0.0006565275,0.11203757,0.04090563,0.00954283,0.024638632,0.80698806],"study_design_scores_gemma":[0.000048550162,0.00012462086,0.0005870325,0.000015629355,0.00006239864,0.000099399345,0.00010141445,0.96418834,0.020043004,0.005937758,0.008773861,0.000017876862],"about_ca_topic_score_codex":0.0030931924,"about_ca_topic_score_gemma":0.0061341804,"teacher_disagreement_score":0.0038779727,"about_ca_system_score_codex":0.0005123839,"about_ca_system_score_gemma":0.0010452509,"threshold_uncertainty_score":0.01297307},"labels":[],"label_agreement":null},{"id":"W4391567383","doi":"10.1007/978-3-031-45190-4_17","title":"Machine Learning, Features, and Computational Approaches to Discourse Analysis","year":2023,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut Universitaire de Gériatrie de Montréal","funders":"","keywords":"Computer science; Artificial intelligence","score_opus":0.08122349840416718,"score_gpt":0.2611066355419474,"score_spread":0.17988313713778023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391567383","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006662077,0.046750586,0.71303153,0.0063186567,0.0012456503,0.00012436342,0.0015805833,0.0025799817,0.2217066],"genre_scores_gemma":[0.22672819,0.043280885,0.45501423,0.0009850498,0.004215825,0.0005144912,0.0048801536,0.0017651052,0.262616],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9995807,0.00015833988,0.000021660637,0.00008467544,0.0001361956,0.000018384246],"domain_scores_gemma":[0.99790096,0.0017874534,0.000057008852,0.00010926285,0.000111826535,0.000033520297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007901208,0.00080702396,0.00070616667,0.002391476,0.00062207045,0.003767951,0.00088599604,0.00070227217,0.015189391],"category_scores_gemma":[0.00400555,0.00033400464,0.0005319334,0.0040811626,0.0012557188,0.004660951,0.00063192996,0.0018142891,0.0048415232],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029815868,0.000041491745,0.00039743193,0.000362269,0.000021129936,0.00004770509,0.0005965119,0.0024154931,0.001193788,0.4723528,0.066696465,0.45584515],"study_design_scores_gemma":[0.000004867659,0.00001748357,0.0009547416,0.00019359792,0.000018587818,0.0002148264,0.00034136168,0.029958643,0.0010004538,0.73024756,0.23702714,0.000020664313],"about_ca_topic_score_codex":0.0015051257,"about_ca_topic_score_gemma":0.0021157688,"teacher_disagreement_score":0.015189391,"about_ca_system_score_codex":0.001131461,"about_ca_system_score_gemma":0.0006769526,"threshold_uncertainty_score":0.050813556},"labels":[],"label_agreement":null},{"id":"W4391572936","doi":"10.1007/s10618-023-01001-y","title":"VEM$$^2$$L: an easy but effective framework for fusing text and structure knowledge on sparse knowledge graph completion","year":2024,"lang":"en","type":"article","venue":"Data Mining and Knowledge Discovery","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute","funders":"National Natural Science Foundation of China","keywords":"Computer science; Knowledge graph; Graph; Information retrieval; Artificial intelligence; Theoretical computer science","score_opus":0.05877690887851249,"score_gpt":0.33299489273397936,"score_spread":0.27421798385546686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391572936","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013113485,0.000094512536,0.99481213,0.00020207449,0.00004806249,0.00006950744,0.00037006117,0.002507496,0.000584938],"genre_scores_gemma":[0.060909376,0.00016935865,0.9314928,0.00032448012,0.00012583734,0.0002848813,0.0020433706,0.0006495243,0.0040002964],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99770385,0.00087282114,0.00011203799,0.00049551,0.0006216493,0.00019408116],"domain_scores_gemma":[0.9953257,0.0023658464,0.0001862434,0.0011783816,0.00067507726,0.000268701],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028021804,0.0010837709,0.0015312692,0.0029350636,0.0010534448,0.0020679317,0.004154458,0.0021507242,0.0094571095],"category_scores_gemma":[0.014455513,0.00078081974,0.0017758176,0.0022367097,0.0010913735,0.004477466,0.0058511784,0.0036972503,0.0032320174],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003550678,0.000385514,0.0011661638,0.000431916,0.00014613257,0.00021501145,0.0002872201,0.07831257,0.0058614244,0.12237177,0.056133628,0.73433363],"study_design_scores_gemma":[0.000024502964,0.00004156484,0.00013564053,0.00003083565,0.000015858777,0.000062372594,0.000040697243,0.88224447,0.0020311968,0.10617034,0.009178226,0.000024262969],"about_ca_topic_score_codex":0.0074877464,"about_ca_topic_score_gemma":0.016486667,"teacher_disagreement_score":0.0094571095,"about_ca_system_score_codex":0.0010359155,"about_ca_system_score_gemma":0.0019655004,"threshold_uncertainty_score":0.03163719},"labels":[],"label_agreement":null},{"id":"W4391680357","doi":"10.1145/3640460","title":"Revisiting Bag of Words Document Representations for Efficient Ranking with Transformers","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Canadian Institute of Steel Construction","keywords":"Computer science; Transformer; Search engine indexing; Inference; ENCODE; Artificial intelligence; Question answering; Information retrieval; Machine learning; Natural language processing; Voltage","score_opus":0.01808328787664678,"score_gpt":0.2747271749205496,"score_spread":0.25664388704390284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391680357","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021344563,0.00075238314,0.96886337,0.0004104358,0.00009733298,0.00012148235,0.0007541176,0.005766901,0.0018894015],"genre_scores_gemma":[0.4349084,0.0012432147,0.5524042,0.00039514533,0.00024516584,0.00027580236,0.0036510066,0.0006606335,0.0062163603],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985863,0.00051293534,0.00012963373,0.00022175866,0.0004187786,0.00013055642],"domain_scores_gemma":[0.9970489,0.0013723214,0.00018844001,0.00082921836,0.00048558827,0.00007565438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019057274,0.000965622,0.0014159003,0.0019362109,0.00043848925,0.002566481,0.0019167236,0.0008113269,0.0044460064],"category_scores_gemma":[0.009744365,0.0004510947,0.00084602175,0.0032035888,0.00078233494,0.006673744,0.0015836465,0.001889517,0.0040550525],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000537569,0.0002724984,0.0013435963,0.00044006598,0.00009212266,0.000115712675,0.00027262472,0.14950787,0.015935384,0.07433572,0.01776751,0.7393792],"study_design_scores_gemma":[0.00006751157,0.00017565211,0.00023747377,0.000024597175,0.00003275599,0.00010699955,0.00008598909,0.91654956,0.0071161976,0.07019403,0.005374593,0.00003457674],"about_ca_topic_score_codex":0.005040421,"about_ca_topic_score_gemma":0.0073216767,"teacher_disagreement_score":0.005040421,"about_ca_system_score_codex":0.001008998,"about_ca_system_score_gemma":0.001992778,"threshold_uncertainty_score":0.014873326},"labels":[],"label_agreement":null},{"id":"W4391681217","doi":"10.1145/3643681","title":"VeriGen: A Large Language Model for Verilog Code Generation","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Design Automation of Electronic Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":191,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Army Research Office; National Science Foundation","keywords":"Computer science; Verilog; Correctness; Programming language; Set (abstract data type); Scripting language; Code (set theory); Hardware description language; Compiler; Embedded system; Field-programmable gate array","score_opus":0.04508213604149722,"score_gpt":0.29203474536437435,"score_spread":0.24695260932287713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391681217","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018292518,0.0007986176,0.8030402,0.0008131145,0.00034839366,0.00053374906,0.007907146,0.162097,0.0061691655],"genre_scores_gemma":[0.20286396,0.0007809312,0.73611414,0.0010915959,0.00011008832,0.0012876809,0.028594872,0.02246803,0.0066887587],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99847215,0.0004868041,0.00012979332,0.00034026662,0.00045399656,0.00011701676],"domain_scores_gemma":[0.9949809,0.0028901652,0.00025051576,0.001133195,0.0006451623,0.000100162986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001835361,0.0017518731,0.00048123783,0.0011610388,0.00035927,0.0013536775,0.0035335168,0.0012387678,0.013218997],"category_scores_gemma":[0.011622193,0.0010805267,0.0016987781,0.00065991265,0.0007661986,0.002691297,0.0015620951,0.002556879,0.007100868],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008148248,0.00041456133,0.0057382374,0.00199781,0.00019676007,0.00048273272,0.0003032015,0.5134848,0.02077499,0.025783626,0.11363207,0.31637648],"study_design_scores_gemma":[0.00013230932,0.00013396356,0.00030086,0.00009573999,0.0000391655,0.00014693892,0.00002846925,0.94221276,0.0128105655,0.013173167,0.030894617,0.000031432297],"about_ca_topic_score_codex":0.0041107317,"about_ca_topic_score_gemma":0.008934466,"teacher_disagreement_score":0.013218997,"about_ca_system_score_codex":0.0011537666,"about_ca_system_score_gemma":0.0031564555,"threshold_uncertainty_score":0.044221938},"labels":[],"label_agreement":null},{"id":"W4391682754","doi":"10.1016/j.neuron.2024.01.016","title":"Data science opportunities of large language models for neuroscience and biomedicine","year":2024,"lang":"en","type":"article","venue":"Neuron","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute; Montreal Neurological Institute and Hospital","funders":"","keywords":"Biomedicine; Cognitive reframing; Cognitive science; Cognitive neuroscience; Neuroscience; Computational neuroscience; Cognition; Computer science; Psychology; Data science; Biology; Bioinformatics","score_opus":0.17261567622721516,"score_gpt":0.35244863348445266,"score_spread":0.1798329572572375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391682754","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004100155,0.03014718,0.8547006,0.099520646,0.0013179255,0.000101042046,0.0013624536,0.0011125576,0.007637408],"genre_scores_gemma":[0.1945682,0.047260176,0.72143734,0.015604071,0.01240231,0.001048127,0.0024273747,0.00094805314,0.004304309],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99355894,0.0044163913,0.0002615021,0.0007113811,0.0009183752,0.00013338985],"domain_scores_gemma":[0.9162349,0.07242248,0.0009579288,0.0069983318,0.0022727659,0.001113633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026169773,0.0017051948,0.0017878525,0.0032138207,0.0013015151,0.008556397,0.0029560681,0.004046661,0.004972469],"category_scores_gemma":[0.06853156,0.0011232307,0.002066638,0.0028305762,0.0062873242,0.025180714,0.006817065,0.014893051,0.0020282315],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000866743,0.00006836958,0.0011870418,0.00044198416,0.00011913581,0.00011084301,0.0006236254,0.009206221,0.0006330044,0.9090922,0.016676912,0.061754014],"study_design_scores_gemma":[0.000011123908,0.000019853784,0.00013744547,0.00014155376,0.000013555538,0.000050424434,0.00008906865,0.039651774,0.00019898244,0.9353618,0.024295947,0.000028490267],"about_ca_topic_score_codex":0.0019055832,"about_ca_topic_score_gemma":0.0020321317,"teacher_disagreement_score":0.026169773,"about_ca_system_score_codex":0.0029196632,"about_ca_system_score_gemma":0.0028186296,"threshold_uncertainty_score":0.13840067},"labels":[],"label_agreement":null},{"id":"W4391689861","doi":"10.1016/j.softx.2024.101649","title":"APRCOIE: An open information extraction system for Chinese","year":2024,"lang":"en","type":"article","venue":"SoftwareX","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada; National Office for Philosophy and Social Sciences","keywords":"Computer science; Information extraction; Information retrieval; World Wide Web; Natural language processing","score_opus":0.02400441237466808,"score_gpt":0.314194521563081,"score_spread":0.29019010918841287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391689861","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031969197,0.0015333066,0.5403139,0.0013794785,0.00043551557,0.0021788075,0.13763474,0.24605593,0.038499072],"genre_scores_gemma":[0.11323473,0.0017434197,0.58100253,0.00046989418,0.00029076228,0.0021401371,0.26686773,0.0070823464,0.02716845],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99912125,0.00011783306,0.00015551843,0.00021477554,0.00032668846,0.00006393691],"domain_scores_gemma":[0.9973984,0.0008932847,0.0002869726,0.00039699097,0.0008833505,0.0001410529],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015457798,0.0013359555,0.00072058564,0.0063411808,0.0010155259,0.0014269421,0.0011815194,0.00059006445,0.012981758],"category_scores_gemma":[0.006052715,0.00053748843,0.00091558543,0.0051938575,0.00044884282,0.004146472,0.0019462013,0.0008096386,0.007840032],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050545926,0.00017225642,0.009173991,0.0025716077,0.00021787756,0.0016041786,0.0021106093,0.0033273983,0.022986084,0.024946703,0.377106,0.55527776],"study_design_scores_gemma":[0.00020089494,0.0001639419,0.014963719,0.0002821764,0.00025831166,0.0012865971,0.000664182,0.07772893,0.047333583,0.021733325,0.83510613,0.0002781033],"about_ca_topic_score_codex":0.009312331,"about_ca_topic_score_gemma":0.0113816615,"teacher_disagreement_score":0.012981758,"about_ca_system_score_codex":0.00090396946,"about_ca_system_score_gemma":0.0044848043,"threshold_uncertainty_score":0.04342836},"labels":[],"label_agreement":null},{"id":"W4391815000","doi":"10.1016/j.artmed.2024.102814","title":"Automated image label extraction from radiology reports — A review","year":2024,"lang":"en","type":"review","venue":"Artificial Intelligence in Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"European Regional Development Fund; Fundação para a Ciência e a Tecnologia; Canadian Mennonite University","keywords":"Computer science; Artificial intelligence; Annotation; Variety (cybernetics); Field (mathematics); Natural language processing; Medical imaging; Information retrieval; Information extraction; Machine learning","score_opus":0.16428139144750825,"score_gpt":0.4525903850097035,"score_spread":0.28830899356219525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391815000","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00024276652,0.9973501,0.0011727599,0.0002910535,0.00010433822,0.000033273114,0.000087901564,0.00003511038,0.0006827379],"genre_scores_gemma":[0.0017778997,0.99488294,0.002569387,0.00021149404,0.00011582619,0.000040598836,0.00018898094,0.000012959263,0.00019993114],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979857,0.000521819,0.00045658852,0.0003278916,0.00064859417,0.000059328417],"domain_scores_gemma":[0.987356,0.009336221,0.0011807561,0.00029502957,0.0017137966,0.00011819209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005103582,0.0012886144,0.0022828802,0.011983898,0.00045757776,0.0021820364,0.0023288713,0.0015064784,0.003232858],"category_scores_gemma":[0.0137421265,0.0008143208,0.0022618521,0.008148849,0.0010751432,0.0035690537,0.0011629725,0.0012038895,0.0018682781],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055966204,0.00004662841,0.0005910012,0.07186518,0.00035734422,0.000077440985,0.00013875501,0.00033459772,0.0005184884,0.0013189019,0.012601438,0.9120943],"study_design_scores_gemma":[0.00005100285,0.00020801951,0.005975479,0.11812176,0.0033231091,0.0021141255,0.00047699403,0.000947896,0.0027488274,0.0047704773,0.8611221,0.00014016245],"about_ca_topic_score_codex":0.0033867063,"about_ca_topic_score_gemma":0.004464805,"teacher_disagreement_score":0.011983898,"about_ca_system_score_codex":0.0012159529,"about_ca_system_score_gemma":0.0034406388,"threshold_uncertainty_score":0.026990652},"labels":[],"label_agreement":null},{"id":"W4391833165","doi":"10.1145/3640543.3645200","title":"Why and When LLM-Based Assistants Can Go Wrong: Investigating the Effectiveness of Prompt-Based Interactions for Software Help-Seeking","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Context (archaeology); Task (project management); Domain (mathematical analysis); Software; Baseline (sea); Relevance (law); Perception; Software engineering; Psychology; Engineering; Programming language","score_opus":0.044925306135747876,"score_gpt":0.29267780404446536,"score_spread":0.24775249790871748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391833165","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98207724,0.00023268974,0.014792007,0.00015014583,0.000025306777,0.00029407334,0.000075599295,0.0011405909,0.0012122466],"genre_scores_gemma":[0.9558839,0.00011982495,0.04206684,0.00023836696,0.000023814122,0.00039808685,0.00013296507,0.00014897957,0.0009872826],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98915535,0.007523718,0.00063766673,0.0011348991,0.0011670273,0.00038139042],"domain_scores_gemma":[0.8689678,0.11077641,0.008460506,0.0054005305,0.0043249223,0.0020698304],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01100309,0.0008695642,0.0006817722,0.00088520494,0.0005802698,0.0018801364,0.0012247249,0.0015632996,0.0017599956],"category_scores_gemma":[0.11765205,0.0005500543,0.00041815024,0.00033879545,0.00078410795,0.0028397904,0.0018297632,0.0010109984,0.0009794631],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010918344,0.007074443,0.12892312,0.0064182975,0.0003659888,0.0021273117,0.12196195,0.005658621,0.13930537,0.0018415055,0.0050819903,0.5703231],"study_design_scores_gemma":[0.0035985906,0.063188925,0.4266435,0.0029495906,0.0018602285,0.006611443,0.09075286,0.17068094,0.16419499,0.011670384,0.056445554,0.0014029648],"about_ca_topic_score_codex":0.0008993261,"about_ca_topic_score_gemma":0.0013078836,"teacher_disagreement_score":0.01100309,"about_ca_system_score_codex":0.0005224212,"about_ca_system_score_gemma":0.0009874047,"threshold_uncertainty_score":0.058190584},"labels":[],"label_agreement":null},{"id":"W4391840390","doi":"10.1145/3648471","title":"Utilizing BERT for Information Retrieval: Survey, Applications, Resources, and Challenges","year":2024,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":107,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; Central China Normal University; York University","keywords":"Computer science; Information retrieval; Data science; World Wide Web","score_opus":0.16111223945183448,"score_gpt":0.35412702620949055,"score_spread":0.19301478675765607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391840390","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014422334,0.64282244,0.29456878,0.011037174,0.00063933776,0.00036497743,0.0012091297,0.0042461026,0.030689707],"genre_scores_gemma":[0.13115227,0.615271,0.23176591,0.0031021552,0.002180395,0.0003197783,0.004122044,0.00059574324,0.011490777],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976827,0.0007809812,0.00020096182,0.00033201388,0.00085516967,0.00014811625],"domain_scores_gemma":[0.99391234,0.0037613253,0.0001993129,0.00085134903,0.0011341835,0.0001415655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033640028,0.001359226,0.0015282684,0.0059914687,0.0006727871,0.00341865,0.0032569512,0.0018012782,0.0060409103],"category_scores_gemma":[0.009483993,0.0009903975,0.0010271412,0.00914656,0.0014814035,0.011302468,0.0025048319,0.002166199,0.005461649],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008936492,0.00014107667,0.0011794792,0.0035892155,0.00007456898,0.000059392074,0.00023939845,0.006957298,0.0022843594,0.037783083,0.02870558,0.9188971],"study_design_scores_gemma":[0.0000785443,0.0005423453,0.0033560272,0.0027828056,0.0002579739,0.0016707785,0.0014013987,0.2616963,0.009207585,0.14921023,0.56950706,0.00028892368],"about_ca_topic_score_codex":0.00600151,"about_ca_topic_score_gemma":0.0041372064,"teacher_disagreement_score":0.0060409103,"about_ca_system_score_codex":0.0015016583,"about_ca_system_score_gemma":0.0022439086,"threshold_uncertainty_score":0.020208836},"labels":[],"label_agreement":null},{"id":"W4391899324","doi":"10.25071/2564-2855.36","title":"Is ChatGPT taking over the language classroom?","year":2024,"lang":"en","type":"article","venue":"Working papers in Applied Linguistics and Linguistics at York","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Linguistics; Psychology; Computer science; Mathematics education; Philosophy","score_opus":0.019031228502891715,"score_gpt":0.2599082870151254,"score_spread":0.2408770585122337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391899324","genre_codex":"empirical","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46477264,0.006232481,0.037988734,0.24004875,0.0066048116,0.00027731064,0.00076250534,0.0037752388,0.23953755],"genre_scores_gemma":[0.95688903,0.001069665,0.003037007,0.008498885,0.0009037813,0.00021436575,0.00020776465,0.0006110382,0.028568475],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.98421234,0.008162546,0.00042249996,0.0024680605,0.0026144057,0.002120124],"domain_scores_gemma":[0.9571354,0.016767077,0.0037823226,0.005213212,0.005036788,0.012065274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010745615,0.0005869992,0.00091752986,0.0014178305,0.011000314,0.014494933,0.002345095,0.004042262,0.018606028],"category_scores_gemma":[0.053003084,0.00075929094,0.00054990983,0.001534226,0.009422523,0.01884156,0.0103297755,0.006661749,0.005959693],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005851798,0.000554105,0.07346121,0.000721623,0.00007538626,0.003349702,0.3717172,0.00035402348,0.0039427695,0.09073285,0.12018404,0.33432192],"study_design_scores_gemma":[0.000094300456,0.0004254066,0.042728774,0.0010337987,0.00007465886,0.0022472374,0.31162983,0.00201776,0.003232595,0.05108481,0.5852079,0.00022291804],"about_ca_topic_score_codex":0.0095627,"about_ca_topic_score_gemma":0.012286716,"teacher_disagreement_score":0.018606028,"about_ca_system_score_codex":0.004854523,"about_ca_system_score_gemma":0.0056222994,"threshold_uncertainty_score":0.062243402},"labels":[],"label_agreement":null},{"id":"W4391904037","doi":"10.24251/hicss.2023.072","title":"Introduction to the Minitrack on Text Mining and Analytics","year":2023,"lang":"en","type":"article","venue":"Proceedings of the ... Annual Hawaii International Conference on System Sciences/Proceedings of the Annual Hawaii International Conference on System Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Social media; Computer science; Data science; Big data; World Wide Web; Government (linguistics); Social media analytics; Analytics; Knowledge management; Data mining","score_opus":0.05798968165721539,"score_gpt":0.3003729609921053,"score_spread":0.24238327933488993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391904037","genre_codex":"methods","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022779726,0.15528345,0.35883135,0.11946836,0.22875378,0.0023157594,0.010108101,0.019719047,0.10324222],"genre_scores_gemma":[0.008826557,0.14975998,0.20357367,0.045120195,0.20397267,0.0029410534,0.025254657,0.010411925,0.3501392],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99044424,0.0022865492,0.00091005216,0.0015916496,0.0042779227,0.0004896176],"domain_scores_gemma":[0.95634097,0.020224895,0.0013349333,0.0041885315,0.013059966,0.0048507676],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010123419,0.002956432,0.0028188962,0.010796518,0.0032855524,0.016041562,0.0048840567,0.0050992495,0.08542161],"category_scores_gemma":[0.02580985,0.0017941659,0.0022549725,0.012069485,0.003152771,0.018942682,0.00903802,0.012910648,0.08710582],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000060378497,0.00010919909,0.00026308815,0.00048082054,0.000029475674,0.000060943756,0.00017887255,0.0002686166,0.0009740192,0.005077894,0.77023983,0.2222569],"study_design_scores_gemma":[0.000011922825,0.00008496317,0.0004949781,0.0005777617,0.000011035963,0.00021400292,0.00013412154,0.0009055451,0.00041368016,0.0067757573,0.9903228,0.000053347172],"about_ca_topic_score_codex":0.0029212865,"about_ca_topic_score_gemma":0.004679361,"teacher_disagreement_score":0.08542161,"about_ca_system_score_codex":0.0024918967,"about_ca_system_score_gemma":0.004480557,"threshold_uncertainty_score":0.28576374},"labels":[],"label_agreement":null},{"id":"W4391941699","doi":"10.1080/0142159x.2024.2316223","title":"Twelve tips for Natural Language Processing in medical education program evaluation","year":2024,"lang":"en","type":"article","venue":"Medical Teacher","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Centre for Addiction and Mental Health","funders":"","keywords":"Workflow; Computer science; Preprocessor; Process (computing); Medical education; Data science; Artificial intelligence; Medicine; Programming language","score_opus":0.0370885945004089,"score_gpt":0.40961178499880835,"score_spread":0.37252319049839944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391941699","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009696018,0.009505946,0.55995226,0.3833366,0.0045060534,0.00638133,0.00056288,0.007616173,0.018442735],"genre_scores_gemma":[0.02598221,0.0041975453,0.951167,0.010428075,0.00071515027,0.004351807,0.00024502212,0.00088310486,0.0020300972],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.58863217,0.2867737,0.05924098,0.005610201,0.054283738,0.005459112],"domain_scores_gemma":[0.33336344,0.4979902,0.020096146,0.03560052,0.097953506,0.014996206],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.34693822,0.0032940751,0.0029495347,0.011081196,0.00766438,0.019397305,0.006552052,0.0101138,0.0073259575],"category_scores_gemma":[0.5261205,0.0027997112,0.0029406894,0.0073085222,0.01214062,0.027550653,0.017235627,0.023378178,0.005447766],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042978916,0.0009297427,0.004268557,0.004143019,0.00016562485,0.00054108404,0.015215207,0.0026738439,0.0021223892,0.037262563,0.11779064,0.81445754],"study_design_scores_gemma":[0.00081287144,0.0024878953,0.010810362,0.037771385,0.0005149309,0.002813607,0.0487244,0.018857189,0.008780272,0.37333608,0.49350688,0.0015842032],"about_ca_topic_score_codex":0.0031892662,"about_ca_topic_score_gemma":0.007579497,"teacher_disagreement_score":0.34693822,"about_ca_system_score_codex":0.01095098,"about_ca_system_score_gemma":0.0321515,"threshold_uncertainty_score":0.8053415},"labels":[],"label_agreement":null},{"id":"W4391968356","doi":"10.1142/s0218213024500052","title":"Summary Augmenter: A Text Augmentation Framework to Improve Summarization Quality","year":2024,"lang":"en","type":"article","venue":"International Journal of Artificial Intelligence Tools","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Automatic summarization; Computer science; Quality (philosophy); Information retrieval; Philosophy; Epistemology","score_opus":0.07811631283339375,"score_gpt":0.38591626279305946,"score_spread":0.3077999499596657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391968356","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.084835015,0.003550905,0.85906816,0.0011104038,0.00055073603,0.00042766603,0.0052679856,0.040210675,0.00497839],"genre_scores_gemma":[0.42662787,0.0013025678,0.53998905,0.00049306103,0.00069564854,0.00049845537,0.015367459,0.0013049376,0.0137209175],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994986,0.00014289764,0.000053447995,0.00015670559,0.0001116384,0.00003666934],"domain_scores_gemma":[0.99834895,0.0006075653,0.00019817005,0.00031197435,0.00046373607,0.000069546586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012532108,0.0015323573,0.00080348237,0.0017869769,0.0004400803,0.0009899338,0.0010163345,0.00090864854,0.005077329],"category_scores_gemma":[0.004993861,0.0002993357,0.0008853143,0.0009870436,0.0003315232,0.0026825236,0.00096812204,0.001107579,0.002924669],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070976635,0.0002805787,0.0023783345,0.00063511555,0.00016288437,0.0003255877,0.0005298569,0.040475793,0.06892286,0.0037079048,0.026559075,0.8553122],"study_design_scores_gemma":[0.00018278127,0.0012215769,0.003803499,0.000110624926,0.0003989602,0.0004726179,0.0003605223,0.84801215,0.091882356,0.01170832,0.04174747,0.00009912714],"about_ca_topic_score_codex":0.0019550025,"about_ca_topic_score_gemma":0.0041426676,"teacher_disagreement_score":0.005077329,"about_ca_system_score_codex":0.00039493275,"about_ca_system_score_gemma":0.000821636,"threshold_uncertainty_score":0.016985357},"labels":[],"label_agreement":null},{"id":"W4391973028","doi":"10.1016/j.compbiomed.2024.108189","title":"A comprehensive evaluation of large Language models on benchmark biomedical text processing tasks","year":2024,"lang":"en","type":"article","venue":"Computers in Biology and Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":97,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; York University","keywords":"Benchmark (surveying); Computer science; Task (project management); Set (abstract data type); Domain (mathematical analysis); Work (physics); Artificial intelligence; Engineering; Mathematics","score_opus":0.045177384775807475,"score_gpt":0.377081677449967,"score_spread":0.33190429267415955,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391973028","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5585918,0.080045395,0.23562607,0.00830177,0.0047525163,0.0025495682,0.030756198,0.057804886,0.021571713],"genre_scores_gemma":[0.650445,0.010357725,0.22326307,0.0031807534,0.0011490005,0.0015418406,0.0984834,0.0018089454,0.009770343],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99351656,0.003088065,0.00067544484,0.0015504091,0.00089099776,0.00027849985],"domain_scores_gemma":[0.9873375,0.0085530775,0.00044280375,0.0014383544,0.0016532163,0.00057501945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011640972,0.003574848,0.0018094229,0.0034109713,0.0013074378,0.0019483198,0.002814957,0.002691403,0.0028928707],"category_scores_gemma":[0.024321271,0.0007013856,0.0022014691,0.0024884383,0.0010220001,0.0041056904,0.002118642,0.0028308574,0.0030996555],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031100898,0.0030887676,0.0119886,0.00607045,0.0027156514,0.00073197566,0.00065555295,0.18512838,0.018938242,0.0022641234,0.08785744,0.67745066],"study_design_scores_gemma":[0.0006985343,0.0038661642,0.011054633,0.0006018034,0.0011291899,0.0008058197,0.0008298782,0.9140688,0.03149708,0.0058838394,0.029271834,0.00029233098],"about_ca_topic_score_codex":0.01315187,"about_ca_topic_score_gemma":0.0198573,"teacher_disagreement_score":0.01315187,"about_ca_system_score_codex":0.002051093,"about_ca_system_score_gemma":0.003300952,"threshold_uncertainty_score":0.06156403},"labels":[],"label_agreement":null},{"id":"W4391986589","doi":"10.1145/3643991.3645074","title":"Can ChatGPT Support Developers? An Empirical Evaluation of Large Language Models for Code Generation","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Programming language; Code generation; Code (set theory); Code review; Software engineering; Static program analysis; Software; Software development; Operating system","score_opus":0.17828140476919138,"score_gpt":0.41025927595688766,"score_spread":0.23197787118769628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391986589","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97156715,0.0006154574,0.018473042,0.00096560665,0.000053898682,0.00045421426,0.0011988336,0.0029635895,0.0037082057],"genre_scores_gemma":[0.96423876,0.00032396338,0.027769145,0.00029573278,0.00003682643,0.00090827193,0.0047067287,0.00057731825,0.0011434181],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98020923,0.015738146,0.0007044063,0.0014265589,0.0015940472,0.0003276728],"domain_scores_gemma":[0.75047374,0.22015373,0.007016737,0.013715268,0.00510544,0.0035351114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021089436,0.00075694476,0.0005448987,0.0015038563,0.00088296336,0.0021512194,0.001459982,0.0014006611,0.0018668649],"category_scores_gemma":[0.16024797,0.00050008245,0.00047859235,0.001228456,0.0011571937,0.004109416,0.0026468427,0.0021213868,0.0011963905],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0052118474,0.0057391096,0.2983019,0.004602653,0.00039349997,0.0012183883,0.071307994,0.032496378,0.014127913,0.00775176,0.033506896,0.52534163],"study_design_scores_gemma":[0.0024136705,0.008988094,0.2801032,0.0019924503,0.00075288356,0.0019070285,0.034035478,0.5187798,0.020838602,0.021275181,0.10839847,0.00051512953],"about_ca_topic_score_codex":0.003269503,"about_ca_topic_score_gemma":0.0048929085,"teacher_disagreement_score":0.021089436,"about_ca_system_score_codex":0.0014691141,"about_ca_system_score_gemma":0.0019898782,"threshold_uncertainty_score":0.11153293},"labels":[],"label_agreement":null},{"id":"W4391992289","doi":"10.1109/eiecs59936.2023.10435469","title":"Triple-Compressed BERT for Efficient Implementation on NLP Tasks","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing","score_opus":0.051941081025909074,"score_gpt":0.33835829713474186,"score_spread":0.2864172161088328,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391992289","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030200321,0.00035469886,0.9300816,0.00050799147,0.00017142271,0.00015215029,0.0016433415,0.03036265,0.0065258113],"genre_scores_gemma":[0.44236606,0.0004522808,0.5445205,0.0003227418,0.000080693535,0.00031838394,0.004687975,0.0014995119,0.00575183],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99956506,0.00007625843,0.00003811705,0.00008017778,0.00018120644,0.000059112113],"domain_scores_gemma":[0.9991215,0.00033125991,0.00004518699,0.00029196835,0.00016699087,0.000043138214],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00062078726,0.0007805065,0.0005484335,0.00066530146,0.00041206685,0.0011547634,0.0017825492,0.00073793536,0.009573799],"category_scores_gemma":[0.00458357,0.0004061065,0.00053188874,0.0010355558,0.00052383036,0.0032758117,0.0015455709,0.0015917199,0.003257773],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010028637,0.00024870367,0.0020139667,0.0004913909,0.000079702935,0.00044868942,0.00055322953,0.2582983,0.039343227,0.06581201,0.030760836,0.6009471],"study_design_scores_gemma":[0.00002871547,0.000047335998,0.00014889686,0.000018973977,0.000012404302,0.0000901808,0.00007440893,0.9605019,0.014103753,0.016284218,0.00867279,0.000016391554],"about_ca_topic_score_codex":0.009501039,"about_ca_topic_score_gemma":0.013912868,"teacher_disagreement_score":0.009573799,"about_ca_system_score_codex":0.0009171968,"about_ca_system_score_gemma":0.0018609406,"threshold_uncertainty_score":0.032027483},"labels":[],"label_agreement":null},{"id":"W4392110993","doi":"10.3390/modelling5010016","title":"Intent Identification by Semantically Analyzing the Search Query","year":2024,"lang":"en","type":"article","venue":"Modelling—International Open Access Journal of Modelling in Engineering Science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Web search query; Computer science; Identification (biology); Information retrieval; Query expansion; Web query classification; Query optimization; Search engine","score_opus":0.06832558378806373,"score_gpt":0.3581409375351658,"score_spread":0.28981535374710204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392110993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1314804,0.0006902203,0.84979606,0.00044553218,0.000057534784,0.00038332518,0.0012463151,0.010540102,0.0053604445],"genre_scores_gemma":[0.744673,0.00040972052,0.24852169,0.00022592978,0.00006499701,0.00016497813,0.0026063733,0.00029545982,0.0030378758],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99910504,0.00021954947,0.00008782828,0.00023645055,0.00025988359,0.000091265414],"domain_scores_gemma":[0.9984648,0.00057212624,0.00020360132,0.0003061961,0.00037732886,0.000075926735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010707078,0.0009800208,0.0010164052,0.0025957355,0.0003856141,0.0013877387,0.00073370605,0.00081868644,0.0016294632],"category_scores_gemma":[0.0045639924,0.00029064814,0.00096252206,0.0013060737,0.00040146534,0.0029745363,0.0017726506,0.0008830038,0.0017897886],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015676667,0.0006454832,0.025607709,0.00083491043,0.00019492491,0.00066732476,0.0029241038,0.015918665,0.2400777,0.0200131,0.009552957,0.6819955],"study_design_scores_gemma":[0.000042296157,0.00044155185,0.009354256,0.0000766184,0.00014419096,0.00094624946,0.0008360153,0.90215385,0.055058256,0.023087304,0.0077512874,0.00010810554],"about_ca_topic_score_codex":0.0027116179,"about_ca_topic_score_gemma":0.0029422154,"teacher_disagreement_score":0.0027116179,"about_ca_system_score_codex":0.00040350287,"about_ca_system_score_gemma":0.0007954323,"threshold_uncertainty_score":0.005662501},"labels":[],"label_agreement":null},{"id":"W4392144222","doi":"10.1111/cogs.13413","title":"Determining the Relativity of Word Meanings Through the Construction of Individualized Models of Semantic Memory","year":2024,"lang":"en","type":"article","venue":"Cognitive Science","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Word (group theory); Computer science; Natural language processing; Semantics (computer science); Linguistics; Cognitive science; Artificial intelligence; Psychology; Philosophy; Programming language","score_opus":0.05578627591767514,"score_gpt":0.30205567213022433,"score_spread":0.24626939621254917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392144222","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53720707,0.0003383829,0.45469916,0.00025999654,0.00002278242,0.0001074021,0.00030031643,0.0003606453,0.0067043183],"genre_scores_gemma":[0.9080551,0.0001725259,0.09078115,0.000021503725,0.0000132104005,0.00015786274,0.00032850972,0.00008239739,0.000387628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974348,0.0012788647,0.00020658359,0.00066891476,0.00030707594,0.00010382782],"domain_scores_gemma":[0.9888524,0.0077437046,0.00097296876,0.0017950981,0.0005083444,0.00012748242],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041820076,0.0004768619,0.0006872362,0.004167111,0.00093683513,0.004437878,0.0011666851,0.0006899955,0.0015962169],"category_scores_gemma":[0.02842449,0.00070856145,0.0012674696,0.0027402772,0.002790291,0.0091371965,0.0028598134,0.0013852286,0.00031672698],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051943044,0.0002791209,0.1265455,0.00062687375,0.0007162705,0.00043843838,0.03621333,0.06223205,0.015668858,0.34867355,0.001844098,0.40624252],"study_design_scores_gemma":[0.000043201817,0.00019395088,0.056672867,0.00015443232,0.00023535544,0.0006180009,0.010498884,0.38407856,0.007328399,0.5317637,0.008242201,0.00017043199],"about_ca_topic_score_codex":0.0026375554,"about_ca_topic_score_gemma":0.0036830537,"teacher_disagreement_score":0.004437878,"about_ca_system_score_codex":0.0013175375,"about_ca_system_score_gemma":0.0008345621,"threshold_uncertainty_score":0.02211684},"labels":[],"label_agreement":null},{"id":"W4392153463","doi":"10.1109/upcon59197.2023.10434380","title":"Neuro-Symbolic AI: Integrating Symbolic Reasoning with Deep Learning","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Horizon College and Seminary","funders":"","keywords":"Computer science; Artificial intelligence; The Symbolic; Symbolic communication; Cognitive science; Psychology; History","score_opus":0.013650972244720585,"score_gpt":0.24251783345820174,"score_spread":0.22886686121348115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392153463","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010790276,0.00041907266,0.97464097,0.001015566,0.00008581576,0.00006892624,0.00012765481,0.003462347,0.00938929],"genre_scores_gemma":[0.42359757,0.00088496314,0.56798816,0.00051611767,0.00008668502,0.00013117638,0.00047395643,0.00026296044,0.0060583544],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99949014,0.00011473446,0.00004041918,0.00010932847,0.00019883677,0.0000465206],"domain_scores_gemma":[0.99916196,0.00029453833,0.000082001214,0.00024866112,0.00013620088,0.000076734585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00096252473,0.0006410714,0.00051123113,0.0007003923,0.0006119174,0.0023926813,0.00220232,0.00081097253,0.003812275],"category_scores_gemma":[0.0027913384,0.00037820172,0.0006491275,0.00073512265,0.0014246114,0.0033096576,0.0028344442,0.0018481049,0.0009895901],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001948471,0.00023592649,0.0025381655,0.00034570738,0.00022551997,0.0003279766,0.00091457,0.20939715,0.024821397,0.22654265,0.00852302,0.525933],"study_design_scores_gemma":[0.000011267567,0.000041118394,0.00025390123,0.00005469169,0.000035760477,0.00008626091,0.00007473494,0.84748334,0.0073253307,0.13152994,0.0130784195,0.000025199828],"about_ca_topic_score_codex":0.0054472773,"about_ca_topic_score_gemma":0.008337825,"teacher_disagreement_score":0.0054472773,"about_ca_system_score_codex":0.0010324134,"about_ca_system_score_gemma":0.0017009111,"threshold_uncertainty_score":0.012753367},"labels":[],"label_agreement":null},{"id":"W4392193048","doi":"10.1038/s41591-024-02855-5","title":"Adapted large language models can outperform medical experts in clinical text summarization","year":2024,"lang":"en","type":"article","venue":"Nature Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":664,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"National Institute of Biomedical Imaging and Bioengineering; Foundation for the National Institutes of Health; National Institute of Arthritis and Musculoskeletal and Skin Diseases; Agency for Healthcare Research and Quality; National Institutes of Health; National Heart, Lung, and Blood Institute","keywords":"Automatic summarization; Computer science; Natural language processing; Artificial intelligence; Information retrieval; Medicine","score_opus":0.022096485540582234,"score_gpt":0.3445173154923743,"score_spread":0.32242082995179205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392193048","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21120128,0.013157015,0.7369706,0.0060000187,0.0016735167,0.0006040697,0.007028454,0.017214356,0.0061506405],"genre_scores_gemma":[0.7999325,0.002700444,0.1680123,0.0017300425,0.0017573916,0.00036738644,0.017473662,0.000920566,0.0071057887],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972128,0.0016175811,0.00024700747,0.00053251337,0.00027320487,0.00011694899],"domain_scores_gemma":[0.9832591,0.013792264,0.0005079204,0.00070238824,0.001445,0.00029342598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056297663,0.0015414666,0.001116136,0.002364814,0.0005290118,0.0017946443,0.0010038333,0.0018353132,0.002782171],"category_scores_gemma":[0.022216372,0.0005136384,0.0014437346,0.0012980143,0.00030368744,0.0021927091,0.0011077225,0.0021588411,0.0032708754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037124185,0.0008791186,0.010627248,0.0013802412,0.0015668352,0.00046122857,0.0007401545,0.15228611,0.024951335,0.0017756094,0.05325057,0.7483692],"study_design_scores_gemma":[0.00023820224,0.00048847456,0.0035553006,0.00008113742,0.00064064417,0.00017143619,0.00017450964,0.9717281,0.008998472,0.007135444,0.006720969,0.000067268804],"about_ca_topic_score_codex":0.0030542733,"about_ca_topic_score_gemma":0.0063584503,"teacher_disagreement_score":0.0056297663,"about_ca_system_score_codex":0.0005169972,"about_ca_system_score_gemma":0.0013572049,"threshold_uncertainty_score":0.029773414},"labels":[],"label_agreement":null},{"id":"W4392231600","doi":"10.1007/s10044-024-01213-y","title":"Big topic modeling based on a two-level hierarchical latent Beta-Liouville allocation for large-scale data and parameter streaming","year":2024,"lang":"en","type":"article","venue":"Pattern Analysis and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Scale (ratio); BETA (programming language); Big data; Artificial intelligence; Data mining; Physics","score_opus":0.07047955069751044,"score_gpt":0.3082657168572723,"score_spread":0.23778616615976184,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392231600","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005623547,0.00019706582,0.993414,0.00014225382,0.000034459586,0.000027183296,0.00006658165,0.00028888148,0.00020610509],"genre_scores_gemma":[0.37475154,0.0010169449,0.6122284,0.00053424994,0.0006122773,0.00087715255,0.0019345939,0.00066390017,0.0073809824],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99728346,0.0014020345,0.00014656025,0.00055770704,0.00034683087,0.00026344816],"domain_scores_gemma":[0.9891951,0.007915907,0.00041216848,0.001180571,0.0008287808,0.000467512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00614347,0.0010583174,0.0027580145,0.0017412887,0.0015305891,0.0024302271,0.004156327,0.0023556566,0.0035840024],"category_scores_gemma":[0.017231671,0.0013036453,0.0025319506,0.002747174,0.0017372759,0.0048478334,0.0032271992,0.0039675017,0.0013786972],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046407746,0.0003171176,0.0036115323,0.0003183425,0.00033616155,0.00028486346,0.000830774,0.60829854,0.005709331,0.2221396,0.008772423,0.14891724],"study_design_scores_gemma":[0.0000069398775,0.000007743418,0.0000797248,0.0000032979146,0.000008218305,0.000009775518,0.000009688988,0.9810464,0.00013221768,0.018434187,0.00025363275,0.000008209632],"about_ca_topic_score_codex":0.0077492497,"about_ca_topic_score_gemma":0.012840015,"teacher_disagreement_score":0.0077492497,"about_ca_system_score_codex":0.0015718514,"about_ca_system_score_gemma":0.0022228695,"threshold_uncertainty_score":0.032490134},"labels":[],"label_agreement":null},{"id":"W4392305414","doi":"10.5220/0012351200003654","title":"Information Retrieval Chatbot on Military Policies and Standards","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Government of Canada; Simon Fraser University; Department of National Defence","funders":"","keywords":"Chatbot; Computer science; Information retrieval; World Wide Web","score_opus":0.012314301875751523,"score_gpt":0.2629558411921978,"score_spread":0.2506415393164463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392305414","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08069157,0.0019465375,0.4710018,0.02799733,0.0030166497,0.0008995224,0.018457659,0.13094044,0.2650486],"genre_scores_gemma":[0.61429566,0.0011225124,0.125605,0.0058335927,0.0017474528,0.00083928264,0.02785572,0.012396355,0.2103044],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9972318,0.0015183034,0.00011950903,0.0002509188,0.00060232484,0.00027714725],"domain_scores_gemma":[0.9859483,0.010281181,0.00030062738,0.0015013254,0.0011158467,0.0008527155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00390232,0.000814376,0.0009602907,0.0018792534,0.0019373179,0.002452188,0.0011756277,0.0026230472,0.056972194],"category_scores_gemma":[0.016661366,0.00057140255,0.0005863958,0.001406132,0.00081218936,0.0057661943,0.0028992223,0.002688819,0.021372866],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013867574,0.000656003,0.0022086715,0.0005587878,0.000055948494,0.00058622233,0.002040359,0.00905808,0.008599362,0.07857621,0.736421,0.15985256],"study_design_scores_gemma":[0.0004817478,0.00047467268,0.0038926885,0.00032023087,0.000075737364,0.00049282,0.0019832947,0.22344981,0.018291555,0.09803768,0.6523288,0.00017091377],"about_ca_topic_score_codex":0.005556845,"about_ca_topic_score_gemma":0.0065286267,"teacher_disagreement_score":0.056972194,"about_ca_system_score_codex":0.0017256046,"about_ca_system_score_gemma":0.0014142974,"threshold_uncertainty_score":0.19059098},"labels":[],"label_agreement":null},{"id":"W4392353984","doi":"10.18280/ria.380122","title":"Hybrid Approach for Automated Answer Scoring Using Semantic Analysis in Long Hindi Text","year":2024,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Hindi; Natural language processing; Computer science; Artificial intelligence; Information retrieval; Linguistics; Philosophy","score_opus":0.05962336056621872,"score_gpt":0.3058849300047682,"score_spread":0.24626156943854946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392353984","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054022633,0.00051886484,0.9300926,0.0003765757,0.00011325934,0.0003949296,0.0007020207,0.00951074,0.004268389],"genre_scores_gemma":[0.35719168,0.00021113691,0.6330189,0.00013696685,0.000081097634,0.00024603895,0.0015819395,0.0002065459,0.007325608],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982931,0.000670025,0.00013919281,0.00038592078,0.000421191,0.00009055532],"domain_scores_gemma":[0.99812776,0.00067383284,0.00015425136,0.0001951462,0.00077283755,0.0000762524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001773075,0.0008655458,0.0007994194,0.0029127183,0.00052653253,0.0013560376,0.00084599643,0.00078415504,0.003920402],"category_scores_gemma":[0.003388767,0.00023997911,0.0007846708,0.0013225129,0.00036881177,0.0019898487,0.0009440792,0.00072664605,0.0029569706],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027805706,0.0002699075,0.0025354896,0.00025483288,0.00009747009,0.00015648632,0.0004907079,0.008520648,0.02929555,0.004257512,0.005364264,0.948479],"study_design_scores_gemma":[0.00006363289,0.00047625095,0.006997619,0.00005067134,0.00009519579,0.00039469262,0.0008643646,0.9318929,0.03438511,0.011551287,0.013141098,0.000087184286],"about_ca_topic_score_codex":0.002440548,"about_ca_topic_score_gemma":0.004510227,"teacher_disagreement_score":0.003920402,"about_ca_system_score_codex":0.00050369697,"about_ca_system_score_gemma":0.0009830006,"threshold_uncertainty_score":0.013115048},"labels":[],"label_agreement":null},{"id":"W4392397955","doi":"10.61577/jaiar.2024.100005","title":"Reranking passages with coarse-to- ne neural retriever enhanced by list-context information","year":2024,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence and Robotics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Labrador Retriever; Context (archaeology); Computer science; Information retrieval; Neuroscience; Psychology; Communication; Biology; Medicine; Surgery; Paleontology","score_opus":0.02869825168991653,"score_gpt":0.2677436004846747,"score_spread":0.2390453487947582,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392397955","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13934103,0.0031996602,0.8392118,0.00044839233,0.00033591458,0.00045847878,0.000799015,0.011556717,0.0046489784],"genre_scores_gemma":[0.61028206,0.0012692523,0.3680444,0.00046700527,0.0004411762,0.0003042718,0.0027276326,0.0004066496,0.016057495],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948186,0.000095477946,0.00004881655,0.00016335404,0.00014894563,0.0000614967],"domain_scores_gemma":[0.99922657,0.00024059376,0.00008397239,0.00014934716,0.00024293929,0.000056637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007289404,0.0016646335,0.0016982129,0.0018752759,0.00067279994,0.0010328837,0.0016217682,0.0015443894,0.0045935973],"category_scores_gemma":[0.0029339192,0.0004051016,0.0008545142,0.0014922034,0.00043649122,0.002801333,0.0010184214,0.0012637549,0.0019629025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069459167,0.00055060926,0.0020555751,0.000476648,0.0001640928,0.00046957564,0.0002993148,0.08071489,0.054936826,0.0039354097,0.014154436,0.84154814],"study_design_scores_gemma":[0.00007446826,0.00033499295,0.001218833,0.000018734168,0.000121781406,0.0002144959,0.00008991895,0.9690234,0.020840937,0.0041587856,0.003851848,0.000051834846],"about_ca_topic_score_codex":0.010741499,"about_ca_topic_score_gemma":0.018708386,"teacher_disagreement_score":0.010741499,"about_ca_system_score_codex":0.0006810037,"about_ca_system_score_gemma":0.0009908236,"threshold_uncertainty_score":0.021357954},"labels":[],"label_agreement":null},{"id":"W4392405447","doi":"10.1109/tifs.2024.3372809","title":"(Security) Assertions by Large Language Models","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Information Forensics and Security","topic":"Topic Modeling","field":"Computer Science","cited_by":70,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Office of Naval Research; Intel Corporation","keywords":"Computer science; Programming language; Natural language processing","score_opus":0.00911134536940814,"score_gpt":0.23101535780912044,"score_spread":0.2219040124397123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392405447","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012220568,0.00023199653,0.9772988,0.00074789085,0.00006130342,0.00031444227,0.0011561966,0.004717401,0.0032515025],"genre_scores_gemma":[0.20949616,0.00047337014,0.7809645,0.00039107585,0.00011358565,0.00095038913,0.0030853704,0.001325519,0.0032000442],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9883795,0.0055551897,0.00082618615,0.0011871081,0.003664903,0.00038706523],"domain_scores_gemma":[0.9561854,0.030739216,0.0025768701,0.006948658,0.0032238662,0.00032597515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008669131,0.0016052054,0.0005853594,0.0020185893,0.00078298163,0.003795202,0.0024789835,0.0014591904,0.005678999],"category_scores_gemma":[0.04247289,0.0011674984,0.0027657314,0.0012187323,0.002132152,0.0067197545,0.0033637388,0.00264258,0.0015658857],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003601798,0.00030502863,0.0040933415,0.0007879883,0.00014755779,0.0006493682,0.001977353,0.31011432,0.009341063,0.5622144,0.013317616,0.09669181],"study_design_scores_gemma":[0.0000947432,0.000111162546,0.00028259293,0.00020119897,0.00006883314,0.0002285698,0.00020240062,0.7328943,0.008105324,0.2279138,0.02985405,0.000043006297],"about_ca_topic_score_codex":0.004542458,"about_ca_topic_score_gemma":0.008139155,"teacher_disagreement_score":0.008669131,"about_ca_system_score_codex":0.001786074,"about_ca_system_score_gemma":0.0026416967,"threshold_uncertainty_score":0.045847356},"labels":[],"label_agreement":null},{"id":"W4392413980","doi":"10.2196/49997","title":"A Case Demonstration of the Open Health Natural Language Processing Toolkit From the National COVID-19 Cohort Collaborative and the Researching COVID to Enhance Recovery Programs for a Natural Language Processing System for COVID-19 or Postacute Sequelae of SARS CoV-2 Infection: Algorithm Development and Validation","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; U.S. National Library of Medicine; National Heart, Lung, and Blood Institute","keywords":"Coronavirus disease 2019 (COVID-19); Artificial intelligence; Natural language processing; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Preprint; Task (project management); 2019-20 coronavirus outbreak; Computer science; Medicine; World Wide Web; Pathology; Disease; Engineering","score_opus":0.05779017490292085,"score_gpt":0.42539791881473193,"score_spread":0.36760774391181106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392413980","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3752586,0.0022193175,0.4537077,0.08168645,0.0018656508,0.0032945836,0.013340255,0.009543852,0.05908361],"genre_scores_gemma":[0.5808748,0.0014873986,0.37722474,0.009902256,0.00044682642,0.0021448005,0.009365924,0.002275302,0.016278004],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9900303,0.0054326705,0.0007864719,0.0011418337,0.0021483994,0.00046027527],"domain_scores_gemma":[0.95988536,0.032156996,0.00093997683,0.0024366872,0.0030122562,0.001568606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01655025,0.00081280177,0.0005081783,0.0014715756,0.0035659263,0.0024852564,0.00197009,0.0034958853,0.005605928],"category_scores_gemma":[0.04206096,0.0004362928,0.0009814407,0.0010183961,0.001986955,0.00248065,0.0033814188,0.0034759175,0.0019571856],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001697095,0.0010604127,0.08510508,0.0026112113,0.00020176424,0.14031322,0.08648294,0.019725116,0.016162187,0.046709742,0.26909253,0.33083862],"study_design_scores_gemma":[0.00043903515,0.00078489527,0.036489606,0.001975499,0.00020168506,0.07289511,0.03322059,0.13479339,0.03470657,0.043988742,0.6400258,0.00047920417],"about_ca_topic_score_codex":0.022157904,"about_ca_topic_score_gemma":0.0390189,"teacher_disagreement_score":0.022157904,"about_ca_system_score_codex":0.0032528949,"about_ca_system_score_gemma":0.0049086,"threshold_uncertainty_score":0.087527156},"labels":[],"label_agreement":null},{"id":"W4392425838","doi":"10.48550/arxiv.2403.00126","title":"FAC$^2$E: Better Understanding Large Language Model Capabilities by Dissociating Language and Cognition","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Cognition; Psychology; Cognitive psychology; Computer science; Cognitive science; Linguistics; Natural language processing; Philosophy; Neuroscience","score_opus":0.061674440794839504,"score_gpt":0.20253384979864666,"score_spread":0.14085940900380717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392425838","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0962551,0.0005145722,0.8782799,0.0011108342,0.000070062975,0.00029952463,0.0013910446,0.010618348,0.011460545],"genre_scores_gemma":[0.5973754,0.00023776172,0.3961288,0.00027773657,0.000042474443,0.00024992274,0.002333247,0.0006115693,0.0027431115],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99817765,0.0008206005,0.00014303645,0.00032983473,0.00040536295,0.00012354892],"domain_scores_gemma":[0.9870595,0.007773398,0.00070291024,0.0031843176,0.0009514911,0.00032836178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004297504,0.0014919704,0.0006288826,0.002221229,0.0005538392,0.0038855397,0.0012829499,0.0012812568,0.007728222],"category_scores_gemma":[0.023028553,0.0004256248,0.0010337303,0.0011104663,0.0009011329,0.008042537,0.0039175483,0.0018600923,0.0017347514],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007010744,0.00057787116,0.03163393,0.0006418541,0.0003850469,0.00017118009,0.0015840345,0.09031087,0.017498646,0.06452893,0.013920746,0.77804583],"study_design_scores_gemma":[0.00005556108,0.00028942083,0.010568496,0.00007985931,0.00011959388,0.00022183459,0.0005628031,0.84435016,0.024277152,0.108511366,0.010832992,0.00013085346],"about_ca_topic_score_codex":0.0067827133,"about_ca_topic_score_gemma":0.009608395,"teacher_disagreement_score":0.007728222,"about_ca_system_score_codex":0.0010848339,"about_ca_system_score_gemma":0.0019296977,"threshold_uncertainty_score":0.025853455},"labels":[],"label_agreement":null},{"id":"W4392514011","doi":"10.21203/rs.3.rs-4006730/v1","title":"TriDeepRec: A Hybrid Deep Learning Approach to Content and Behaviour-based Recommendation Systems","year":2024,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Content (measure theory); Computer science; Recommender system; Artificial intelligence; Deep learning; Machine learning; Mathematics","score_opus":0.17749923930080155,"score_gpt":0.37935794672821327,"score_spread":0.20185870742741172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392514011","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030349826,0.0008539683,0.95742005,0.00051148265,0.00008611837,0.00018207385,0.00069213496,0.0070508956,0.0028534564],"genre_scores_gemma":[0.48977625,0.0005718469,0.49666893,0.0007677258,0.0000730246,0.00029200633,0.0018771698,0.00017574248,0.009797347],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994122,0.00010240611,0.0000496684,0.0002006581,0.00016620933,0.000068888134],"domain_scores_gemma":[0.9990828,0.00030750735,0.00007149238,0.00015720412,0.00032566465,0.000055343953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011440978,0.0007370011,0.00074679754,0.0008178393,0.0003414078,0.00084490847,0.0022553317,0.0014356866,0.0024729478],"category_scores_gemma":[0.0026172055,0.0006461026,0.0007048854,0.0009331448,0.00034821115,0.0013896333,0.0011206745,0.0020960434,0.0009728776],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027856603,0.00063305843,0.0050066942,0.00021620956,0.00042807515,0.00017718103,0.00013975677,0.47812557,0.012960042,0.0052965065,0.010825431,0.48591292],"study_design_scores_gemma":[0.000008346375,0.000031237578,0.00031075213,0.000006891188,0.000011428024,0.000016722328,0.000004615886,0.9964228,0.0014081196,0.00094848505,0.00082312425,0.0000073851884],"about_ca_topic_score_codex":0.024388839,"about_ca_topic_score_gemma":0.03911385,"teacher_disagreement_score":0.024388839,"about_ca_system_score_codex":0.0012284186,"about_ca_system_score_gemma":0.0011351213,"threshold_uncertainty_score":0.048493803},"labels":[],"label_agreement":null},{"id":"W4392560653","doi":"10.1145/3649449","title":"Pre-Trained Language Models for Text Generation: A Survey","year":2024,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":174,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Natural Science Foundation of China","keywords":"Computer science; Text generation; Key (lock); Artificial intelligence; Field (mathematics); Language model; Deep learning; Data science; Natural language processing; Machine learning","score_opus":0.16776380191581788,"score_gpt":0.3821377924828532,"score_spread":0.2143739905670353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392560653","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008073562,0.57444453,0.38682944,0.0042654984,0.0013256953,0.00025330042,0.0009818526,0.0035706523,0.020255428],"genre_scores_gemma":[0.096093506,0.67721885,0.19313747,0.0018963143,0.002470643,0.00052031106,0.0050601037,0.0010742014,0.022528633],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99938285,0.00018561084,0.000061931234,0.00014614713,0.00019199048,0.000031371965],"domain_scores_gemma":[0.9968194,0.002428487,0.000101809404,0.00021156385,0.00038729174,0.000051420262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013880242,0.0012431601,0.0011466017,0.0018697681,0.00026216108,0.001146818,0.0020817143,0.0011684991,0.0067103407],"category_scores_gemma":[0.006355435,0.0005322486,0.0010012963,0.0020677603,0.00050645124,0.003007324,0.0008140408,0.0015395778,0.0053875763],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005434263,0.000100778176,0.00065305346,0.0028139292,0.00006801456,0.000070385904,0.00013500871,0.01482582,0.0014551005,0.009354644,0.025560325,0.9449087],"study_design_scores_gemma":[0.000071725466,0.0004608629,0.0032097215,0.0042088474,0.00033036596,0.0012243849,0.00048815546,0.24958825,0.013215664,0.04258508,0.68445504,0.0001618206],"about_ca_topic_score_codex":0.0028364111,"about_ca_topic_score_gemma":0.0029916782,"teacher_disagreement_score":0.0067103407,"about_ca_system_score_codex":0.0007216529,"about_ca_system_score_gemma":0.0013482089,"threshold_uncertainty_score":0.02244836},"labels":[],"label_agreement":null},{"id":"W4392608406","doi":"10.33137/ijournal.v9i1.42237","title":"Conversational Breakdown Detector for a Motivational Interviewing Conversational Agent","year":2023,"lang":"en","type":"article","venue":"The iJournal Student Journal of the Faculty of Information","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Interview; Motivational interviewing; Psychology; Detector; Computer science; Sociology; Telecommunications","score_opus":0.04879673560062339,"score_gpt":0.3038524488520255,"score_spread":0.2550557132514021,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392608406","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1886762,0.0016615269,0.78645176,0.0012632125,0.00032207416,0.0009633592,0.002613196,0.0074952794,0.010553409],"genre_scores_gemma":[0.6584377,0.00024273174,0.33133367,0.00032789088,0.00011262793,0.00058447744,0.0036570067,0.00018737241,0.005116467],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99885917,0.0003468085,0.000055703033,0.00039537184,0.00020925813,0.00013369779],"domain_scores_gemma":[0.9986229,0.0006016182,0.00015999193,0.00010768378,0.00036545732,0.00014234959],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018840439,0.00078495836,0.00051612395,0.0015893074,0.0007808473,0.0010070269,0.00088554923,0.0010668194,0.0017528681],"category_scores_gemma":[0.004632945,0.00027674568,0.00056048075,0.00047255235,0.00044724252,0.000931029,0.0014396097,0.0016192456,0.0010966419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018066877,0.00069918035,0.0482475,0.0010137064,0.0002952421,0.0009908328,0.0060633556,0.034116887,0.089105375,0.01610595,0.040083084,0.7614722],"study_design_scores_gemma":[0.000047038957,0.00032990018,0.020712327,0.00010577641,0.0001265372,0.00059643964,0.0016466923,0.9019384,0.028547391,0.011267113,0.034607846,0.000074502605],"about_ca_topic_score_codex":0.0047833533,"about_ca_topic_score_gemma":0.006495359,"teacher_disagreement_score":0.0047833533,"about_ca_system_score_codex":0.0009138928,"about_ca_system_score_gemma":0.0010129629,"threshold_uncertainty_score":0.00996387},"labels":[],"label_agreement":null},{"id":"W4392637304","doi":"10.18653/v1/2023.eval4nlp-1.3","title":"Delving into Evaluation Metrics for Generation: A Thorough Assessment of How Metrics Generalize to Rephrasing Across Languages","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University","keywords":"Computer science; Generative grammar; Heuristics; Natural language processing; Artificial intelligence; Variety (cybernetics); Natural language generation; Phrase; Robustness (evolution); Natural language","score_opus":0.17340541247336058,"score_gpt":0.4540214667908573,"score_spread":0.2806160543174967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392637304","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3044323,0.012072009,0.65427667,0.0026858984,0.0007971685,0.0018771706,0.0032933324,0.0058998973,0.014665585],"genre_scores_gemma":[0.7143228,0.0015434793,0.27477026,0.00048553536,0.00017569111,0.0012594225,0.0039698165,0.0021558884,0.0013170438],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.90430266,0.054550905,0.010302591,0.0059441333,0.023145353,0.0017543875],"domain_scores_gemma":[0.6107544,0.28758296,0.022846546,0.03518556,0.0410734,0.0025572465],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08420991,0.002582379,0.002062141,0.008602167,0.0013121563,0.006432684,0.0024847055,0.0023777063,0.0011567611],"category_scores_gemma":[0.34403703,0.0005881226,0.0016609771,0.007994606,0.0026722294,0.009500186,0.0045956234,0.0028116154,0.00066857645],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015522653,0.00084832567,0.11408406,0.004013508,0.001978109,0.00031363455,0.007209231,0.092627995,0.021096261,0.031514514,0.012936438,0.71182567],"study_design_scores_gemma":[0.000248159,0.008218689,0.14334957,0.0037656548,0.0014712921,0.0020849975,0.008977581,0.59166527,0.07199585,0.10462563,0.062238496,0.001358764],"about_ca_topic_score_codex":0.0042204987,"about_ca_topic_score_gemma":0.003487182,"teacher_disagreement_score":0.9157901,"about_ca_system_score_codex":0.0028264446,"about_ca_system_score_gemma":0.0021633643,"threshold_uncertainty_score":0.44534987},"labels":[],"label_agreement":null},{"id":"W4392669708","doi":"10.18653/v1/2023.ijcnlp-main.35","title":"Attacking Open-domain Question Answering by Injecting Misinformation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"National Science Foundation","keywords":"Misinformation; Chen; Computer science; Domain (mathematical analysis); Question answering; Computational linguistics; Open domain; Natural language processing; Artificial intelligence; Linguistics; Library science; Philosophy; Mathematics; Geology; Computer security","score_opus":0.03142707791162644,"score_gpt":0.3068473172510588,"score_spread":0.2754202393394324,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1114957,0.004233007,0.8292815,0.020198364,0.001437367,0.0005430674,0.0014162506,0.014520569,0.016874237],"genre_scores_gemma":[0.8481891,0.000809074,0.1390285,0.0044360366,0.0007291262,0.00020803286,0.0015588381,0.00055079476,0.0044905134],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9822726,0.010472178,0.00074415037,0.0018812495,0.0038010564,0.00082881737],"domain_scores_gemma":[0.9203648,0.059242405,0.0022870766,0.012728448,0.004350514,0.0010269011],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014406716,0.0012427564,0.00202216,0.00205825,0.0020219714,0.0035364106,0.0023636804,0.005088559,0.0034222875],"category_scores_gemma":[0.06888351,0.0010404984,0.00130138,0.0017236056,0.0027162395,0.011806348,0.009269629,0.0057389466,0.0027197383],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003468874,0.001427917,0.014032775,0.0012090061,0.0009664129,0.0020192685,0.00706737,0.06699548,0.048564065,0.16157436,0.13045363,0.5622209],"study_design_scores_gemma":[0.00015519932,0.00021543325,0.0008812359,0.00012515532,0.00024300178,0.000724368,0.0010281138,0.711615,0.033167135,0.22343627,0.028310394,0.00009873464],"about_ca_topic_score_codex":0.0019983992,"about_ca_topic_score_gemma":0.0019162646,"teacher_disagreement_score":0.014406716,"about_ca_system_score_codex":0.0011040884,"about_ca_system_score_gemma":0.0016240684,"threshold_uncertainty_score":0.07619089},"labels":[],"label_agreement":null},{"id":"W4392669730","doi":"10.18653/v1/2023.eval4nlp-1.16","title":"Reference-Free Summarization Evaluation with Large Language Models","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; National Research Council Canada","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Natural language processing; Multi-document summarization; Artificial intelligence","score_opus":0.06768538163856708,"score_gpt":0.2967039691196833,"score_spread":0.2290185874811162,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669730","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2286416,0.0032705113,0.68813527,0.00065243006,0.0006873858,0.0009909949,0.004416373,0.06612419,0.007081217],"genre_scores_gemma":[0.6700061,0.00043846195,0.30271605,0.00027950841,0.0002084702,0.00077758793,0.018578617,0.002712889,0.0042823227],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9851664,0.009841145,0.00094602606,0.0020360972,0.0016749031,0.0003354595],"domain_scores_gemma":[0.970303,0.018082464,0.0010492216,0.003798056,0.00602075,0.0007465422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011490072,0.0026964429,0.0016010185,0.002266319,0.00092843524,0.0021926432,0.0020025212,0.0022163729,0.004413959],"category_scores_gemma":[0.039737996,0.00050287065,0.001126296,0.0015706313,0.00072041596,0.003973198,0.002500739,0.0018925609,0.003164494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003843812,0.0012835417,0.0049397647,0.0027286776,0.00089794106,0.0008152232,0.0025985586,0.22610995,0.07528121,0.0071200957,0.043959316,0.6304219],"study_design_scores_gemma":[0.00042426033,0.0019143849,0.002898298,0.000092350914,0.00024313157,0.00024808582,0.00062423164,0.9224698,0.052445352,0.0064475574,0.01206098,0.00013166369],"about_ca_topic_score_codex":0.0030441657,"about_ca_topic_score_gemma":0.0038863162,"teacher_disagreement_score":0.011490072,"about_ca_system_score_codex":0.0012031776,"about_ca_system_score_gemma":0.0013532477,"threshold_uncertainty_score":0.06076604},"labels":[],"label_agreement":null},{"id":"W4392669732","doi":"10.18653/v1/2023.ijcnlp-main.71","title":"PACT: Pretraining with Adversarial Contrastive Learning for Text Classification","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Adversarial system; Computer science; Linguistics; Natural language processing; Artificial intelligence; Joint (building); Islam; Computational linguistics; Contrastive analysis; Philosophy; Engineering; Theology","score_opus":0.047870754172629246,"score_gpt":0.27268264358549654,"score_spread":0.22481188941286728,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669732","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0105551435,0.002122448,0.9635733,0.00095396244,0.0008753312,0.00022399942,0.00073543517,0.015986944,0.004973604],"genre_scores_gemma":[0.30030254,0.0017382102,0.65033376,0.0024839772,0.0011268302,0.0013567961,0.008287591,0.0026956904,0.031674653],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990758,0.00032972905,0.00004679796,0.00026082597,0.00018097185,0.00010578094],"domain_scores_gemma":[0.99810874,0.0011550783,0.00006400359,0.0002894255,0.00028890208,0.000093885756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021387802,0.0020412246,0.0013231314,0.0010382873,0.00079137465,0.0014613047,0.003162557,0.0022684354,0.009765064],"category_scores_gemma":[0.00533701,0.0009755953,0.0013335661,0.0010289963,0.0008764041,0.002963512,0.0027824156,0.0060458514,0.008263043],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004466583,0.00035151685,0.0010538709,0.00024803326,0.0002540603,0.00024302937,0.0001463321,0.2326393,0.006648668,0.011213247,0.09407151,0.6526838],"study_design_scores_gemma":[0.000023210041,0.000055807224,0.000119240736,0.000018089017,0.00001372944,0.000035082063,0.00001404006,0.9872104,0.0022491321,0.0072674644,0.0029828073,0.00001099981],"about_ca_topic_score_codex":0.0050095334,"about_ca_topic_score_gemma":0.008601931,"teacher_disagreement_score":0.009765064,"about_ca_system_score_codex":0.00085903436,"about_ca_system_score_gemma":0.0013114442,"threshold_uncertainty_score":0.0326674},"labels":[],"label_agreement":null},{"id":"W4392669746","doi":"10.18653/v1/2023.ijcnlp-srw.12","title":"Evaluating Large Language Models’ Understanding of Financial Terminology via Definition Modeling","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institute of Advanced Industrial Science and Technology; Nvidia","keywords":"Terminology; Computer science; Computational linguistics; Linguistics; Natural language processing; Philosophy","score_opus":0.2625755630681034,"score_gpt":0.3487654642913454,"score_spread":0.086189901223242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669746","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37950775,0.0044813114,0.5842782,0.004993004,0.0004167935,0.00053392857,0.0076194895,0.005881647,0.01228779],"genre_scores_gemma":[0.7971631,0.00083865784,0.17935586,0.00045186182,0.00016199001,0.0003958038,0.019367779,0.0007809565,0.0014839339],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98797613,0.008611389,0.00069594244,0.0014771784,0.0010056329,0.00023370552],"domain_scores_gemma":[0.9458712,0.04773547,0.0010540212,0.0024110775,0.0022573038,0.000670961],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01225194,0.0016650789,0.0012535063,0.0039569763,0.0012560281,0.0063271616,0.0025774448,0.001983283,0.0033817384],"category_scores_gemma":[0.048113286,0.0010374053,0.0027208105,0.0029374436,0.000894553,0.012801536,0.003531301,0.0034394423,0.0013959563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029466122,0.0012782111,0.032821573,0.0015146711,0.002118145,0.0009319363,0.0061461064,0.3480887,0.0088274935,0.080236636,0.058596447,0.45649347],"study_design_scores_gemma":[0.00013275254,0.00011911068,0.0023371,0.000066027926,0.00022132398,0.000111684414,0.0010281506,0.9487824,0.0025947096,0.03866382,0.0058913785,0.000051457722],"about_ca_topic_score_codex":0.008999459,"about_ca_topic_score_gemma":0.013069179,"teacher_disagreement_score":0.01225194,"about_ca_system_score_codex":0.002335088,"about_ca_system_score_gemma":0.0018277812,"threshold_uncertainty_score":0.064795256},"labels":[],"label_agreement":null},{"id":"W4392669853","doi":"10.18653/v1/2023.ijcnlp-main.46","title":"Analyzing and Predicting Persistence of News Tweets","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Persistence (discontinuity); Computational linguistics; Computer science; Joint (building); Volume (thermodynamics); Association (psychology); Natural language processing; Artificial intelligence; Linguistics; Data science; Library science; Engineering; Philosophy; Epistemology","score_opus":0.05018962379470674,"score_gpt":0.25700656681081985,"score_spread":0.20681694301611311,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669853","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9694559,0.0041742185,0.011070925,0.0013207041,0.00045705953,0.0000685882,0.009836595,0.00061575946,0.0030001537],"genre_scores_gemma":[0.9792838,0.0011071342,0.0044212025,0.00007147475,0.00057553715,0.000050399587,0.0126359435,0.00007147735,0.0017829174],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99959606,0.00008416229,0.000031971253,0.00011939797,0.000092452,0.00007602632],"domain_scores_gemma":[0.995937,0.0023725613,0.0005428529,0.00027374722,0.0005942208,0.00027971453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009955748,0.0004586238,0.00046196987,0.0039718905,0.0005085582,0.0013078736,0.00038364346,0.0005343146,0.0008370708],"category_scores_gemma":[0.006534713,0.00027856854,0.00040191432,0.0024412933,0.00019865832,0.0015341142,0.0006335588,0.00079354417,0.0012031685],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017510483,0.00052208186,0.6543933,0.0006714941,0.0006025894,0.00066098524,0.001152,0.009761646,0.01538663,0.0026016773,0.05060002,0.2618965],"study_design_scores_gemma":[0.00010043856,0.000524319,0.5023532,0.00020076596,0.0007329815,0.0007560747,0.0035721234,0.4408259,0.013718831,0.006519729,0.03057699,0.00011852187],"about_ca_topic_score_codex":0.004078113,"about_ca_topic_score_gemma":0.006790382,"teacher_disagreement_score":0.004078113,"about_ca_system_score_codex":0.00032168423,"about_ca_system_score_gemma":0.0002803975,"threshold_uncertainty_score":0.008108795},"labels":[],"label_agreement":null},{"id":"W4392669868","doi":"10.18653/v1/2023.findings-ijcnlp.16","title":"The Glass Ceiling of Automatic Evaluation in Natural Language Generation","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Grand Équipement National De Calcul Intensif","keywords":"Computer science; Readability; Metric (unit); Fidelity; Field (mathematics); Rank (graph theory); Machine learning; Data mining; Natural language; Natural language generation; Artificial intelligence; Data science; Programming language; Engineering","score_opus":0.04116435527696538,"score_gpt":0.3149507560117089,"score_spread":0.27378640073474353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669868","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035870656,0.026874674,0.82000184,0.049259588,0.0042646946,0.00054910156,0.0017412051,0.027381847,0.034056343],"genre_scores_gemma":[0.57429564,0.0052182027,0.36061144,0.009677776,0.005834432,0.00081177487,0.0048925923,0.008359581,0.030298574],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9352949,0.04312854,0.0020789353,0.0057812263,0.0115205,0.0021957995],"domain_scores_gemma":[0.8211507,0.14277034,0.0014075668,0.015079242,0.016261006,0.003331204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.055514712,0.0023297798,0.004121942,0.0031832962,0.0021488962,0.01096073,0.003570341,0.005753794,0.022765128],"category_scores_gemma":[0.1253439,0.001789884,0.0015699017,0.0019966548,0.005842511,0.01735499,0.0073725046,0.008901615,0.011757254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015460256,0.0004823072,0.004135788,0.0008337683,0.00024032191,0.00012887621,0.00069776416,0.010916148,0.007206712,0.05043031,0.16163163,0.76175034],"study_design_scores_gemma":[0.00030500386,0.0009020455,0.0039645177,0.00065363845,0.00017037036,0.00038139155,0.0007670926,0.60929734,0.019129246,0.2812136,0.08299794,0.00021782868],"about_ca_topic_score_codex":0.007055924,"about_ca_topic_score_gemma":0.0063225226,"teacher_disagreement_score":0.055514712,"about_ca_system_score_codex":0.003381253,"about_ca_system_score_gemma":0.0033517405,"threshold_uncertainty_score":0.29359335},"labels":[],"label_agreement":null},{"id":"W4392669874","doi":"10.18653/v1/2023.ijcnlp-main.39","title":"ProMap: Effective Bilingual Lexicon Induction via Language Model Prompting","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; University of British Columbia","funders":"","keywords":"Lexicon; Computer science; Linguistics; Computational linguistics; Natural language processing; Artificial intelligence; Natural language; Philosophy","score_opus":0.024252014903650625,"score_gpt":0.2873487763397928,"score_spread":0.26309676143614213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669874","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014946873,0.00070638506,0.78782487,0.000710254,0.0007405733,0.00041817874,0.0049252645,0.17805243,0.0116751725],"genre_scores_gemma":[0.17875576,0.0005520219,0.76861274,0.0006078046,0.00037806775,0.00090716017,0.02823806,0.0117749395,0.010173494],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978235,0.00088521955,0.00015501202,0.00059355976,0.00034819648,0.00019456628],"domain_scores_gemma":[0.99673355,0.001566333,0.000120094686,0.000793118,0.000595278,0.00019166768],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017972157,0.0019469201,0.001660174,0.0021360372,0.001411908,0.0025294432,0.0030030124,0.0012224057,0.025478901],"category_scores_gemma":[0.005915307,0.0013491971,0.0013716976,0.0022057404,0.00073014037,0.005919921,0.008582,0.0024964213,0.026201375],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010108694,0.00042720485,0.0011451499,0.0009599783,0.00013765316,0.00086002133,0.0005698023,0.0051233326,0.034256972,0.017910464,0.15735224,0.7802464],"study_design_scores_gemma":[0.00152761,0.0006229196,0.0017496306,0.00019251912,0.00032212026,0.0020891784,0.0013432023,0.52807325,0.13219938,0.1441765,0.18743883,0.00026489736],"about_ca_topic_score_codex":0.0021559033,"about_ca_topic_score_gemma":0.005125064,"teacher_disagreement_score":0.025478901,"about_ca_system_score_codex":0.0007066719,"about_ca_system_score_gemma":0.002838574,"threshold_uncertainty_score":0.08523542},"labels":[],"label_agreement":null},{"id":"W4392669908","doi":"10.18653/v1/2023.ijcnlp-main.31","title":"Interactive-Chain-Prompting: Ambiguity Resolution for Crosslingual Conditional Generation with Interaction","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Ambiguity; Computer science; Computational linguistics; Natural language processing; Artificial intelligence; Linguistics; Philosophy; Programming language","score_opus":0.08344911601284492,"score_gpt":0.338691458608126,"score_spread":0.25524234259528106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669908","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004380963,0.00022336614,0.9271511,0.00028373548,0.00023977265,0.00017229577,0.0005820575,0.06326435,0.0037024263],"genre_scores_gemma":[0.177275,0.00022752411,0.8052209,0.00040886988,0.00026727584,0.00047459482,0.0022015567,0.007831313,0.006093],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976947,0.0011695011,0.00014309588,0.000502526,0.00030496393,0.00018524006],"domain_scores_gemma":[0.9927289,0.005158624,0.00017732476,0.0010978518,0.00056361477,0.00027371928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042173816,0.0024579419,0.0012966909,0.0014256381,0.0012590669,0.0025363837,0.003856926,0.0022421074,0.05153406],"category_scores_gemma":[0.014654659,0.001267309,0.0011864037,0.0010971847,0.0012143638,0.0047211535,0.007531093,0.0026226444,0.01293954],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024360674,0.00042658707,0.001510781,0.0009940207,0.00017068813,0.0009421329,0.0027086877,0.018317305,0.024898684,0.06469975,0.10616162,0.7767337],"study_design_scores_gemma":[0.00050142483,0.00019845234,0.00048681354,0.00014535169,0.00013102384,0.00049879926,0.0005148037,0.693717,0.047800858,0.1842816,0.07155719,0.00016667738],"about_ca_topic_score_codex":0.0020540827,"about_ca_topic_score_gemma":0.0032971497,"teacher_disagreement_score":0.05153406,"about_ca_system_score_codex":0.00073986204,"about_ca_system_score_gemma":0.0013651308,"threshold_uncertainty_score":0.17239857},"labels":[],"label_agreement":null},{"id":"W4392669913","doi":"10.18653/v1/2023.findings-ijcnlp.6","title":"Learning to Diversify Neural Text Generation via Degenerative Model","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Language model; Limiting; Artificial intelligence; Diversity (politics); Repetition (rhetorical device); Machine learning; Natural language processing; Engineering; Linguistics","score_opus":0.06359297354936129,"score_gpt":0.27038437903035406,"score_spread":0.20679140548099278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669913","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18969163,0.00075114914,0.802829,0.0008616054,0.000094975745,0.00017315803,0.00012855169,0.0028078712,0.0026620533],"genre_scores_gemma":[0.88414407,0.00022375314,0.111248404,0.00053934415,0.000102469996,0.00037056525,0.00041678952,0.00016242459,0.0027921922],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909043,0.0003264956,0.000058548707,0.00029385148,0.0001377664,0.000092778755],"domain_scores_gemma":[0.993497,0.0042631994,0.0005193501,0.0007395166,0.0007239333,0.0002569396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028782072,0.0014162961,0.0011268014,0.001144021,0.00047622868,0.0010319548,0.0021749125,0.0018648258,0.0012358375],"category_scores_gemma":[0.012732157,0.00075098174,0.0006797393,0.00063750223,0.000982147,0.0028445798,0.0024073904,0.0024743055,0.0005837056],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036484143,0.0004390525,0.00544267,0.00017757462,0.00013023567,0.00025043695,0.00050373276,0.6994869,0.014679097,0.009277545,0.0041911392,0.26505676],"study_design_scores_gemma":[0.000017028911,0.0000524685,0.0001115697,0.000005737715,0.000011051785,0.000025217338,0.000011785383,0.99456704,0.001426243,0.003549702,0.00021625751,0.000005857882],"about_ca_topic_score_codex":0.0013837538,"about_ca_topic_score_gemma":0.0027178158,"teacher_disagreement_score":0.0028782072,"about_ca_system_score_codex":0.00094228686,"about_ca_system_score_gemma":0.0008168093,"threshold_uncertainty_score":0.015221596},"labels":[],"label_agreement":null},{"id":"W4392669927","doi":"10.18653/v1/2023.findings-ijcnlp.4","title":"PRiSM: Enhancing Low-Resource Document-Level Relation Extraction with Relation-Aware Score Calibration","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"National Supercomputing Center, Korea Institute of Science and Technology Information; Institute for Information and Communications Technology Promotion; Samsung; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Relation (database); Computer science; Calibration; Prism; Relationship extraction; Resource (disambiguation); Code (set theory); Data mining; Key (lock); Source code; Information retrieval; Perspective (graphical); Artificial intelligence; Machine learning; Statistics; Optics; Mathematics; Programming language","score_opus":0.029685417689416287,"score_gpt":0.2561387947892135,"score_spread":0.22645337709979718,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669927","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03947224,0.0044226586,0.86855054,0.0012244172,0.00037540362,0.00037198025,0.009944356,0.06750155,0.008136932],"genre_scores_gemma":[0.3601965,0.0018150156,0.56961364,0.0011010049,0.00058311847,0.0006339284,0.045111094,0.0033677816,0.017577957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964126,0.0009466842,0.00023309227,0.0014429046,0.00074465637,0.00022012233],"domain_scores_gemma":[0.99399877,0.0030563094,0.00038908303,0.0015672048,0.00083088485,0.00015767867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004245141,0.0024935508,0.0014617132,0.0046751257,0.00094088644,0.0027153767,0.0031738682,0.0019311238,0.0054648877],"category_scores_gemma":[0.015996337,0.0007724893,0.0019933907,0.0054981117,0.0008381415,0.0071977866,0.0035581666,0.0036542863,0.009396877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036351455,0.00044570732,0.013028955,0.0008343403,0.0003663225,0.0003350626,0.00081732066,0.039285883,0.015379277,0.010491088,0.09216459,0.82648796],"study_design_scores_gemma":[0.00012828741,0.00019999577,0.0088173505,0.00017834653,0.0002641964,0.0008559629,0.0004029575,0.8613813,0.019205213,0.04237537,0.06605177,0.00013936068],"about_ca_topic_score_codex":0.007690489,"about_ca_topic_score_gemma":0.016760254,"teacher_disagreement_score":0.007690489,"about_ca_system_score_codex":0.0011977485,"about_ca_system_score_gemma":0.0019001086,"threshold_uncertainty_score":0.022450745},"labels":[],"label_agreement":null},{"id":"W4392669929","doi":"10.18653/v1/2023.findings-ijcnlp.30","title":"Mixing It Up: Inducing Empathy and Politeness using Multiple Behaviour-aware Generators for Conversational Systems","year":2023,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Politeness; Computer science; Politeness theory; Encoder; Human–computer interaction; Generator (circuit theory); Empathy; Context (archaeology); Natural language processing; Artificial intelligence; Speech recognition; Psychology; Linguistics; Social psychology","score_opus":0.09577259153113851,"score_gpt":0.29944273510058167,"score_spread":0.20367014356944316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392669929","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057076942,0.0007086338,0.9323256,0.00044839986,0.00014113518,0.0002581392,0.00034432873,0.0048586107,0.0038382472],"genre_scores_gemma":[0.74792725,0.00030361617,0.24206953,0.00037555807,0.00014541567,0.00041141585,0.0016771781,0.0006267109,0.0064633135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985317,0.0008243088,0.000047997775,0.0003530779,0.00013871197,0.0001042959],"domain_scores_gemma":[0.99768853,0.0016672729,0.000102829676,0.00025929793,0.00018244707,0.00009957312],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015839405,0.0009926226,0.000523895,0.00061023363,0.00038898224,0.00095734914,0.0008750862,0.0010716977,0.0029014342],"category_scores_gemma":[0.0062714624,0.00044920732,0.00083263626,0.00029178147,0.0006195525,0.001357927,0.0014535721,0.0014112506,0.0014230317],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009752189,0.00062859897,0.0054493495,0.00065275806,0.0002429083,0.00048136135,0.0023042609,0.11528484,0.073986016,0.017983817,0.014117893,0.76789296],"study_design_scores_gemma":[0.00005995631,0.00015685477,0.0012239045,0.000027361764,0.00005161228,0.00017754441,0.00023071446,0.96207124,0.013963512,0.017026572,0.0049772463,0.0000334074],"about_ca_topic_score_codex":0.0013682799,"about_ca_topic_score_gemma":0.002347162,"teacher_disagreement_score":0.0029014342,"about_ca_system_score_codex":0.0006260509,"about_ca_system_score_gemma":0.00056845124,"threshold_uncertainty_score":0.009706259},"labels":[],"label_agreement":null},{"id":"W4392827844","doi":"10.1007/978-3-031-56066-8_9","title":"Towards Automated End-to-End Health Misinformation Free Search with a Large Language Model","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Misinformation; Computer science; Artificial intelligence; Computer security","score_opus":0.020958864599865398,"score_gpt":0.28650093249714126,"score_spread":0.26554206789727586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392827844","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013642707,0.0005395857,0.95827377,0.001075467,0.000083887,0.00021203047,0.002830159,0.02081716,0.0025251901],"genre_scores_gemma":[0.31950057,0.000398263,0.66004115,0.00081226666,0.00021698703,0.00031033013,0.008159254,0.001410424,0.009150751],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99712974,0.0011476639,0.00021007854,0.00060021627,0.0006624573,0.00024988517],"domain_scores_gemma":[0.9896199,0.0073371157,0.00027762965,0.0015685993,0.0008762768,0.00032049525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030347838,0.0013238966,0.002672115,0.0023506093,0.0011343078,0.0036139044,0.002649514,0.0028490615,0.008106446],"category_scores_gemma":[0.012420462,0.0008812156,0.0019356912,0.0019139715,0.0008710739,0.0059258873,0.0048465338,0.002498366,0.0061510033],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026034983,0.0012355628,0.0043962263,0.00090432743,0.00058512826,0.00091429264,0.0007814864,0.16779138,0.022257544,0.034495752,0.074906245,0.6891286],"study_design_scores_gemma":[0.000046535995,0.00005151743,0.00017651108,0.000013482618,0.00004225569,0.00010005351,0.000081249585,0.96852934,0.0029287087,0.026005328,0.002004642,0.000020390515],"about_ca_topic_score_codex":0.0066625793,"about_ca_topic_score_gemma":0.012131747,"teacher_disagreement_score":0.008106446,"about_ca_system_score_codex":0.0011782718,"about_ca_system_score_gemma":0.002827868,"threshold_uncertainty_score":0.027118742},"labels":[],"label_agreement":null},{"id":"W4392828353","doi":"10.48550/arxiv.2403.08763","title":"Simple and Scalable Strategies to Continually Pre-train Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Oak Ridge National Laboratory; Natural Sciences and Engineering Research Council of Canada; Office of Science; Fonds de recherche du Québec – Nature et technologies; Université de Montréal; Deutscher Akademischer Austauschdienst; Canadian Institute for Advanced Research; Canada Excellence Research Chairs, Government of Canada; U.S. Department of Energy","keywords":"Simple (philosophy); Computer science; Scalability; Epistemology; Philosophy","score_opus":0.0430427905600358,"score_gpt":0.21012120165286874,"score_spread":0.16707841109283295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392828353","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08293797,0.0017035351,0.83522534,0.0012679897,0.00057499245,0.0007578361,0.001445697,0.068940796,0.0071458276],"genre_scores_gemma":[0.4876392,0.0006899245,0.49024543,0.0011621206,0.00021909106,0.001300713,0.005660597,0.0042883595,0.008794479],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99830747,0.00041767902,0.00013785918,0.00069266,0.00031718393,0.00012721214],"domain_scores_gemma":[0.9954541,0.001351414,0.00023551394,0.0020779592,0.0006800156,0.00020099057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029445663,0.0030123133,0.0013940292,0.00095607986,0.0007960837,0.001606141,0.0056694266,0.0015211651,0.00793799],"category_scores_gemma":[0.015991325,0.0015471617,0.0012053474,0.0009263566,0.001134203,0.006653386,0.0038086874,0.0050764256,0.011842731],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081908907,0.0010454208,0.007057068,0.00056265277,0.00039484966,0.00030963184,0.0006335125,0.18202718,0.038179487,0.007249309,0.037818838,0.723903],"study_design_scores_gemma":[0.00021712207,0.00030206947,0.0015866072,0.00005800895,0.00009335944,0.000210517,0.00021638969,0.9456108,0.025974026,0.011364319,0.014255261,0.00011149403],"about_ca_topic_score_codex":0.0082749035,"about_ca_topic_score_gemma":0.018632092,"teacher_disagreement_score":0.0082749035,"about_ca_system_score_codex":0.0011457546,"about_ca_system_score_gemma":0.0019962338,"threshold_uncertainty_score":0.02655524},"labels":[],"label_agreement":null},{"id":"W4392846357","doi":"10.1007/978-3-031-56060-6_26","title":"Adapting Standard Retrieval Benchmarks to Evaluate Generated Answers","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Artificial intelligence","score_opus":0.026618923285977928,"score_gpt":0.2743444726809662,"score_spread":0.24772554939498825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392846357","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39845017,0.009690199,0.38400134,0.0013323713,0.0022641812,0.0027197758,0.019343736,0.13753748,0.04466077],"genre_scores_gemma":[0.5506876,0.0014417217,0.35230342,0.00068563817,0.00047286364,0.0012299211,0.06919033,0.0058931815,0.018095288],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9907488,0.0027707496,0.0010891694,0.0012390082,0.0037190442,0.00043322518],"domain_scores_gemma":[0.9716196,0.014349187,0.0008237833,0.0041526235,0.008336349,0.0007184905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007389854,0.0020624413,0.0016570458,0.0065827584,0.00073254283,0.0036956102,0.0042531085,0.0025167896,0.007951305],"category_scores_gemma":[0.044285823,0.0005740457,0.00089836767,0.004596435,0.00051880715,0.003910356,0.002480269,0.00145234,0.0076988484],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026223876,0.0016655626,0.0073630055,0.0020779415,0.00045350782,0.00031948832,0.00038554377,0.04328428,0.03042448,0.0045526414,0.08880453,0.8180467],"study_design_scores_gemma":[0.0006173987,0.0016594493,0.0058002286,0.00023462836,0.00032289381,0.00043752405,0.0004651037,0.8896874,0.061570466,0.009607523,0.02946205,0.0001354057],"about_ca_topic_score_codex":0.0057744454,"about_ca_topic_score_gemma":0.006376032,"teacher_disagreement_score":0.007951305,"about_ca_system_score_codex":0.0016796077,"about_ca_system_score_gemma":0.0012166145,"threshold_uncertainty_score":0.039081752},"labels":[],"label_agreement":null},{"id":"W4392902227","doi":"10.32920/25418206.v1","title":"Exploration and Mitigation of Stereotypical Gender Biases in Information Retrieval Systems","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Relevance (law); Judgement; Ranking (information retrieval); Debiasing; Computer science; Information retrieval; Set (abstract data type); Artificial intelligence; Machine learning; Psychology; Social psychology","score_opus":0.08431150949639821,"score_gpt":0.28530859426595423,"score_spread":0.20099708476955602,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392902227","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58630955,0.0025111826,0.40190834,0.0015035818,0.0001085282,0.00034076063,0.00032967422,0.0018337816,0.0051546586],"genre_scores_gemma":[0.9101809,0.00037154838,0.08725049,0.00029344534,0.00007380385,0.00014535939,0.00030170652,0.00008911541,0.001293582],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9887412,0.0064115375,0.00079320575,0.0010740756,0.0024996786,0.00048021527],"domain_scores_gemma":[0.98139566,0.009570766,0.002035773,0.0040737344,0.0026301818,0.00029384322],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011282961,0.00072474603,0.00090154505,0.0013481419,0.00067676004,0.0017511077,0.0013898655,0.0010348428,0.0011346807],"category_scores_gemma":[0.04399433,0.00037130932,0.00069156394,0.0010458386,0.0010219972,0.003238745,0.0020761697,0.0009050494,0.0006859368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013022443,0.0005255134,0.030709643,0.00083946757,0.00033307562,0.00023988202,0.0022461077,0.06548654,0.095196344,0.012172727,0.0030162055,0.78793234],"study_design_scores_gemma":[0.00021603836,0.0019535709,0.030781714,0.00018529898,0.0002755084,0.00070896535,0.0014465017,0.74589777,0.14914176,0.05654333,0.012642645,0.00020689346],"about_ca_topic_score_codex":0.0018790639,"about_ca_topic_score_gemma":0.0025369967,"teacher_disagreement_score":0.011282961,"about_ca_system_score_codex":0.0011489274,"about_ca_system_score_gemma":0.0015592589,"threshold_uncertainty_score":0.059670746},"labels":[],"label_agreement":null},{"id":"W4392902241","doi":"10.32920/25418206","title":"Exploration and Mitigation of Stereotypical Gender Biases in Information Retrieval Systems","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Relevance (law); Judgement; Ranking (information retrieval); Debiasing; Computer science; Information retrieval; Set (abstract data type); Premise; Artificial intelligence; Machine learning; Psychology; Social psychology","score_opus":0.08431150949639821,"score_gpt":0.28530859426595423,"score_spread":0.20099708476955602,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392902241","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58630955,0.0025111826,0.40190834,0.0015035818,0.0001085282,0.00034076063,0.00032967422,0.0018337816,0.0051546586],"genre_scores_gemma":[0.9101809,0.00037154838,0.08725049,0.00029344534,0.00007380385,0.00014535939,0.00030170652,0.00008911541,0.001293582],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9887412,0.0064115375,0.00079320575,0.0010740756,0.0024996786,0.00048021527],"domain_scores_gemma":[0.98139566,0.009570766,0.002035773,0.0040737344,0.0026301818,0.00029384322],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011282961,0.00072474603,0.00090154505,0.0013481419,0.00067676004,0.0017511077,0.0013898655,0.0010348428,0.0011346807],"category_scores_gemma":[0.04399433,0.00037130932,0.00069156394,0.0010458386,0.0010219972,0.003238745,0.0020761697,0.0009050494,0.0006859368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013022443,0.0005255134,0.030709643,0.00083946757,0.00033307562,0.00023988202,0.0022461077,0.06548654,0.095196344,0.012172727,0.0030162055,0.78793234],"study_design_scores_gemma":[0.00021603836,0.0019535709,0.030781714,0.00018529898,0.0002755084,0.00070896535,0.0014465017,0.74589777,0.14914176,0.05654333,0.012642645,0.00020689346],"about_ca_topic_score_codex":0.0018790639,"about_ca_topic_score_gemma":0.0025369967,"teacher_disagreement_score":0.011282961,"about_ca_system_score_codex":0.0011489274,"about_ca_system_score_gemma":0.0015592589,"threshold_uncertainty_score":0.059670746},"labels":[],"label_agreement":null},{"id":"W4392906878","doi":"10.32920/25412857.v1","title":"Extracting Source Information From News Articles","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Magna International (Canada)","funders":"","keywords":"Computer science; Credibility; Information retrieval; Viewpoints; Relevance (law); Categorization; Information source (mathematics); Source credibility; Set (abstract data type); Software; Precision and recall; Recall; Data science; Attribution; Ranking (information retrieval); Artificial intelligence; Political science; Psychology","score_opus":0.033204134909525236,"score_gpt":0.2547677027536751,"score_spread":0.22156356784414988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392906878","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28447616,0.0062657488,0.57295716,0.0014834054,0.0007305739,0.0031333629,0.08506499,0.02183264,0.024055967],"genre_scores_gemma":[0.30251276,0.0032710058,0.5810374,0.00011642737,0.00081088673,0.0012573629,0.1038041,0.0011945551,0.005995523],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969733,0.00045097168,0.0005005893,0.0006538752,0.0012042943,0.00021702069],"domain_scores_gemma":[0.97461796,0.01177818,0.0026289485,0.0023291237,0.00815673,0.0004889067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034729836,0.001670817,0.0011453299,0.036591068,0.0013381891,0.0038381182,0.0010788265,0.0012890501,0.0028594665],"category_scores_gemma":[0.022896519,0.0008088775,0.0013213508,0.017278634,0.00041875735,0.00345918,0.0016712626,0.0013908419,0.004040821],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062840054,0.0002946167,0.046709362,0.0026375856,0.0002861799,0.0013329184,0.0031014697,0.0031648462,0.028838553,0.0042839507,0.026442442,0.88227963],"study_design_scores_gemma":[0.00029055294,0.00076231157,0.22529702,0.002316046,0.0022509359,0.0048332983,0.01155852,0.24137391,0.15869987,0.043571692,0.30846798,0.0005778991],"about_ca_topic_score_codex":0.003661342,"about_ca_topic_score_gemma":0.004407744,"teacher_disagreement_score":0.036591068,"about_ca_system_score_codex":0.00083851174,"about_ca_system_score_gemma":0.0019227574,"threshold_uncertainty_score":0.018367112},"labels":[],"label_agreement":null},{"id":"W4392906991","doi":"10.32920/25412863.v1","title":"Incremental Text Clustering Algorithm Using Incremental Learning in COVID-19 Research Papers","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Cluster analysis; Word2vec; Computer science; Data mining; Correlation clustering; Artificial intelligence; Document clustering; tf–idf; Machine learning; Embedding","score_opus":0.1361018834047257,"score_gpt":0.3970309097637997,"score_spread":0.260929026359074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392906991","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047653243,0.0030272375,0.9260066,0.0007713593,0.0009273626,0.0010073638,0.0028413339,0.011912949,0.005852585],"genre_scores_gemma":[0.123852365,0.0010709208,0.8556521,0.00026791642,0.0005460769,0.00081202976,0.010426658,0.00059727265,0.006774647],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997147,0.00034238276,0.00031383542,0.001086785,0.00089414226,0.00021585934],"domain_scores_gemma":[0.9951508,0.0010655135,0.0005186345,0.0007643045,0.0021839929,0.00031674295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002379056,0.0013187596,0.0016379646,0.010394826,0.0015016027,0.0034736276,0.002891695,0.0013648135,0.0041331085],"category_scores_gemma":[0.009590814,0.0005928018,0.0018021069,0.00856024,0.00060274603,0.0030177296,0.0024161693,0.0011973754,0.004156512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003246963,0.00033254476,0.005557755,0.00050701195,0.0002435344,0.0003258279,0.00038195396,0.03030377,0.01012392,0.0066835172,0.022621555,0.92259383],"study_design_scores_gemma":[0.00023740761,0.00038654276,0.0049562505,0.00012494215,0.0002868017,0.000871841,0.0005172526,0.88728076,0.02223294,0.03153949,0.05143949,0.00012629744],"about_ca_topic_score_codex":0.0046565593,"about_ca_topic_score_gemma":0.0062370226,"teacher_disagreement_score":0.010394826,"about_ca_system_score_codex":0.001434896,"about_ca_system_score_gemma":0.0036744832,"threshold_uncertainty_score":0.013826609},"labels":[],"label_agreement":null},{"id":"W4392910835","doi":"10.1109/icassp48485.2024.10446325","title":"A Birgat Model for Multi-Intent Spoken Language Understanding with Hierarchical Semantic Frames","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Optech (Canada)","funders":"National Natural Science Foundation of China","keywords":"Computer science; Utterance; Natural language processing; Artificial intelligence; Pointer (user interface); Spoken language; Generalizability theory; Graph; Margin (machine learning); Machine learning; Theoretical computer science","score_opus":0.09607990330988486,"score_gpt":0.30908444610138114,"score_spread":0.21300454279149628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392910835","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02467704,0.00040301762,0.967176,0.00064559333,0.00007750899,0.00008002395,0.0008794449,0.002248795,0.0038125825],"genre_scores_gemma":[0.7845328,0.00043868233,0.194913,0.00056108297,0.000114094735,0.00044250762,0.0034257877,0.00037940324,0.015192685],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99957687,0.000120378245,0.000020024772,0.00016166686,0.00006715033,0.000053886717],"domain_scores_gemma":[0.9992256,0.00044035155,0.000050288407,0.00010974909,0.00012526725,0.00004869297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010092888,0.0007846689,0.00067123753,0.000837352,0.0005129566,0.0012688362,0.0021507926,0.001559979,0.004681018],"category_scores_gemma":[0.0028516594,0.00050823257,0.0009018896,0.0008074062,0.0006695109,0.0030856952,0.0014241678,0.0026147023,0.0019025573],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005442405,0.00028866917,0.0028179144,0.00024939672,0.00011784363,0.00032368503,0.0013890718,0.50811017,0.010202928,0.1813362,0.013762424,0.28085753],"study_design_scores_gemma":[0.000008733648,0.000019992454,0.00011399122,0.000009762034,0.000010467751,0.00001997035,0.00003099644,0.96251965,0.0005408692,0.035366636,0.0013514039,0.000007501397],"about_ca_topic_score_codex":0.013191451,"about_ca_topic_score_gemma":0.026330642,"teacher_disagreement_score":0.013191451,"about_ca_system_score_codex":0.0015484514,"about_ca_system_score_gemma":0.0013873158,"threshold_uncertainty_score":0.026229382},"labels":[],"label_agreement":null},{"id":"W4392927810","doi":"10.32920/25412863","title":"Incremental Text Clustering Algorithm Using Incremental Learning in COVID-19 Research Papers","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Cluster analysis; Word2vec; Computer science; Data mining; Correlation clustering; Artificial intelligence; Document clustering; Machine learning; Embedding","score_opus":0.1361018834047257,"score_gpt":0.3970309097637997,"score_spread":0.260929026359074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392927810","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047653243,0.0030272375,0.9260066,0.0007713593,0.0009273626,0.0010073638,0.0028413339,0.011912949,0.005852585],"genre_scores_gemma":[0.123852365,0.0010709208,0.8556521,0.00026791642,0.0005460769,0.00081202976,0.010426658,0.00059727265,0.006774647],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997147,0.00034238276,0.00031383542,0.001086785,0.00089414226,0.00021585934],"domain_scores_gemma":[0.9951508,0.0010655135,0.0005186345,0.0007643045,0.0021839929,0.00031674295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002379056,0.0013187596,0.0016379646,0.010394826,0.0015016027,0.0034736276,0.002891695,0.0013648135,0.0041331085],"category_scores_gemma":[0.009590814,0.0005928018,0.0018021069,0.00856024,0.00060274603,0.0030177296,0.0024161693,0.0011973754,0.004156512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003246963,0.00033254476,0.005557755,0.00050701195,0.0002435344,0.0003258279,0.00038195396,0.03030377,0.01012392,0.0066835172,0.022621555,0.92259383],"study_design_scores_gemma":[0.00023740761,0.00038654276,0.0049562505,0.00012494215,0.0002868017,0.000871841,0.0005172526,0.88728076,0.02223294,0.03153949,0.05143949,0.00012629744],"about_ca_topic_score_codex":0.0046565593,"about_ca_topic_score_gemma":0.0062370226,"teacher_disagreement_score":0.010394826,"about_ca_system_score_codex":0.001434896,"about_ca_system_score_gemma":0.0036744832,"threshold_uncertainty_score":0.013826609},"labels":[],"label_agreement":null},{"id":"W4392927976","doi":"10.32920/25412857","title":"Extracting Source Information From News Articles","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Magna International (Canada)","funders":"","keywords":"Computer science; Credibility; Information retrieval; Viewpoints; Information source (mathematics); Relevance (law); Categorization; Source credibility; Software; Precision and recall; Set (abstract data type); Data science; Recall; Attribution; World Wide Web; Artificial intelligence; Political science; Psychology","score_opus":0.033204134909525236,"score_gpt":0.2547677027536751,"score_spread":0.22156356784414988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392927976","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28447616,0.0062657488,0.57295716,0.0014834054,0.0007305739,0.0031333629,0.08506499,0.02183264,0.024055967],"genre_scores_gemma":[0.30251276,0.0032710058,0.5810374,0.00011642737,0.00081088673,0.0012573629,0.1038041,0.0011945551,0.005995523],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969733,0.00045097168,0.0005005893,0.0006538752,0.0012042943,0.00021702069],"domain_scores_gemma":[0.97461796,0.01177818,0.0026289485,0.0023291237,0.00815673,0.0004889067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034729836,0.001670817,0.0011453299,0.036591068,0.0013381891,0.0038381182,0.0010788265,0.0012890501,0.0028594665],"category_scores_gemma":[0.022896519,0.0008088775,0.0013213508,0.017278634,0.00041875735,0.00345918,0.0016712626,0.0013908419,0.004040821],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062840054,0.0002946167,0.046709362,0.0026375856,0.0002861799,0.0013329184,0.0031014697,0.0031648462,0.028838553,0.0042839507,0.026442442,0.88227963],"study_design_scores_gemma":[0.00029055294,0.00076231157,0.22529702,0.002316046,0.0022509359,0.0048332983,0.01155852,0.24137391,0.15869987,0.043571692,0.30846798,0.0005778991],"about_ca_topic_score_codex":0.003661342,"about_ca_topic_score_gemma":0.004407744,"teacher_disagreement_score":0.036591068,"about_ca_system_score_codex":0.00083851174,"about_ca_system_score_gemma":0.0019227574,"threshold_uncertainty_score":0.018367112},"labels":[],"label_agreement":null},{"id":"W4392949175","doi":"10.3233/idt-230629","title":"SExpSMA-based T5: Serial exponential-slime mould algorithm based T5 model for question answer and distractor generation","year":2024,"lang":"en","type":"article","venue":"Intelligent Decision Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Context (archaeology); Range (aeronautics); Exponential function; Algorithm; Artificial intelligence; Process (computing); Natural language processing; Machine learning; Arithmetic; Mathematics; Engineering; Programming language","score_opus":0.05118028656545992,"score_gpt":0.31051808802110653,"score_spread":0.2593378014556466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392949175","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033266045,0.00067156483,0.9589318,0.00041967374,0.00015001751,0.0002514466,0.00032035707,0.0036842066,0.0023048492],"genre_scores_gemma":[0.45660612,0.0004528271,0.5297109,0.00057840464,0.00013175116,0.00071758404,0.001579944,0.0003613885,0.009861057],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989586,0.0002501478,0.000094530966,0.0003819788,0.00021149483,0.000103256745],"domain_scores_gemma":[0.998252,0.0008965197,0.000102733226,0.00013570281,0.00052995497,0.00008321073],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015424071,0.0010983731,0.00092685095,0.0013195937,0.0008504336,0.0012879138,0.0023854293,0.001897944,0.005572202],"category_scores_gemma":[0.0056475345,0.00048439103,0.0018147894,0.0009802819,0.0006336995,0.0019179262,0.0011557735,0.0021701904,0.002047713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073913514,0.0003304085,0.007472928,0.00028623667,0.0001998719,0.00030012912,0.00046642195,0.32810658,0.01511818,0.013455629,0.009297474,0.62422705],"study_design_scores_gemma":[0.000022205011,0.00009222717,0.0003742532,0.000009680772,0.000022080043,0.00006416201,0.000025128218,0.9907432,0.0033316002,0.0033744907,0.0019282807,0.000012669795],"about_ca_topic_score_codex":0.016265478,"about_ca_topic_score_gemma":0.01765173,"teacher_disagreement_score":0.016265478,"about_ca_system_score_codex":0.0014043722,"about_ca_system_score_gemma":0.0021822688,"threshold_uncertainty_score":0.0323416},"labels":[],"label_agreement":null},{"id":"W4392982066","doi":"10.1109/icaiic60209.2024.10463196","title":"Automated Fact Checking Using A Knowledge Graph-based Model","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Graph; Knowledge graph; Programming language; Theoretical computer science; Artificial intelligence","score_opus":0.07634662200306046,"score_gpt":0.3277555577130238,"score_spread":0.2514089357099633,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392982066","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054492317,0.00028967945,0.9335736,0.002522667,0.00006561322,0.00018464071,0.0016878775,0.0022089821,0.004974675],"genre_scores_gemma":[0.7210829,0.00044269642,0.26923335,0.0004120352,0.00007496759,0.00025782894,0.0031066507,0.00016492452,0.0052246503],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991486,0.00021411819,0.000061215265,0.0003070057,0.00018774364,0.000081320795],"domain_scores_gemma":[0.99461925,0.0038966583,0.0004581087,0.0004340652,0.0004641331,0.00012775401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001173536,0.00060557853,0.0005771679,0.002027197,0.0007150184,0.0022406527,0.0017649997,0.0014328306,0.003572783],"category_scores_gemma":[0.007270593,0.00053148915,0.0016572492,0.0012450336,0.0011061229,0.003935207,0.001356035,0.0018740766,0.00059301674],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015149631,0.00012824505,0.0039048991,0.000097145145,0.0000829802,0.0003387738,0.00027490343,0.86859083,0.0011282865,0.065746516,0.0028877012,0.056668162],"study_design_scores_gemma":[0.00000740285,0.000006428126,0.0001787527,0.000009698973,0.000015219866,0.000017997034,0.000012998157,0.9756665,0.00022656123,0.023097133,0.00075594126,0.000005296525],"about_ca_topic_score_codex":0.0340379,"about_ca_topic_score_gemma":0.036556046,"teacher_disagreement_score":0.0340379,"about_ca_system_score_codex":0.0021852145,"about_ca_system_score_gemma":0.0024433017,"threshold_uncertainty_score":0.067679524},"labels":[],"label_agreement":null},{"id":"W4393073548","doi":"10.1007/978-3-031-56063-7_13","title":"Interactive Topic Tagging in Community Question Answering Platforms","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Question answering; Information retrieval; World Wide Web; Natural language processing; Artificial intelligence","score_opus":0.02227528228155629,"score_gpt":0.27430850949419955,"score_spread":0.25203322721264326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393073548","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038870245,0.003811599,0.8627839,0.0029143859,0.0010390325,0.0005269946,0.0054371804,0.031916313,0.052700378],"genre_scores_gemma":[0.25876185,0.0019958902,0.6305527,0.0008692767,0.0011386499,0.0009275132,0.020628903,0.006092845,0.07903235],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99842393,0.0007322378,0.00006547548,0.00028021415,0.00036579452,0.000132409],"domain_scores_gemma":[0.995379,0.0034261309,0.000089497706,0.0005403166,0.00031450842,0.00025054434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027041247,0.00082393317,0.000679737,0.0021848546,0.0014070998,0.0031396912,0.0016491705,0.0016332935,0.021713702],"category_scores_gemma":[0.0072501847,0.0006070377,0.00077851355,0.002706393,0.00043900037,0.009115385,0.0039265086,0.0018099146,0.011365475],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052486773,0.0004753421,0.0020710377,0.0010197419,0.00012470983,0.0004063351,0.0051327646,0.0069791856,0.026849445,0.07013828,0.22552642,0.6607518],"study_design_scores_gemma":[0.000097362376,0.00019494085,0.0027642206,0.0003337576,0.00012446237,0.0005661034,0.002346345,0.23866661,0.024554793,0.1557299,0.5744716,0.00014986162],"about_ca_topic_score_codex":0.0019104524,"about_ca_topic_score_gemma":0.0035956078,"teacher_disagreement_score":0.021713702,"about_ca_system_score_codex":0.00074119464,"about_ca_system_score_gemma":0.00063246384,"threshold_uncertainty_score":0.072639585},"labels":[],"label_agreement":null},{"id":"W4393106418","doi":"10.1007/978-3-031-54752-2_2","title":"AI-Generated Literature","year":2024,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Art","score_opus":0.019552972234056467,"score_gpt":0.23171943870574183,"score_spread":0.21216646647168536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393106418","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020950271,0.010360802,0.07007603,0.004627854,0.0017639931,0.00010033695,0.0015190699,0.0013498986,0.9081069],"genre_scores_gemma":[0.057661492,0.019562902,0.045888994,0.0018355615,0.0023975237,0.00023446858,0.0072054365,0.0014845158,0.8637291],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994259,0.00016670483,0.0000241648,0.00010498092,0.000253119,0.000025056392],"domain_scores_gemma":[0.9980167,0.001192709,0.0000632223,0.0002366464,0.00038137197,0.000109267916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007866432,0.0006933807,0.00037463522,0.0053003337,0.0009830489,0.005434425,0.0012433822,0.0007356142,0.09011858],"category_scores_gemma":[0.004518068,0.0003627737,0.0005362457,0.005495609,0.0009794342,0.005585331,0.0015986527,0.0014190712,0.03167292],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016641432,0.00003639943,0.00021185372,0.00038026902,0.000016991045,0.00009144604,0.0005164054,0.0014440722,0.00054411363,0.4802765,0.28285465,0.23361072],"study_design_scores_gemma":[0.0000051703973,0.0000065217737,0.0002064672,0.00022955387,0.000010865742,0.00015020692,0.0002428916,0.004171261,0.00047520595,0.17538494,0.81910795,0.000009056745],"about_ca_topic_score_codex":0.0016293948,"about_ca_topic_score_gemma":0.0032482988,"teacher_disagreement_score":0.09011858,"about_ca_system_score_codex":0.0016365016,"about_ca_system_score_gemma":0.0015140675,"threshold_uncertainty_score":0.30147672},"labels":[],"label_agreement":null},{"id":"W4393145623","doi":"10.1609/aaai.v38i21.30433","title":"Improving Faithfulness in Abstractive Text Summarization with EDUs Using BART (Student Abstract)","year":2024,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Natural language processing; Computer science; Psychology; Linguistics; Mathematics education; Artificial intelligence; Philosophy","score_opus":0.06324586238832317,"score_gpt":0.3050327422802676,"score_spread":0.24178687989194445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393145623","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15324856,0.002660886,0.82293665,0.00054575293,0.00037452963,0.0002505062,0.0009916035,0.015625719,0.0033657944],"genre_scores_gemma":[0.54517156,0.00076407596,0.4417507,0.00024993485,0.00034854372,0.00023059006,0.004021495,0.00068160635,0.0067815403],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9992938,0.00021279283,0.000072226016,0.00020600627,0.00015859301,0.00005661942],"domain_scores_gemma":[0.99734604,0.0012683464,0.0002885668,0.00032999678,0.00065373175,0.00011339562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012133925,0.0009946589,0.0007864179,0.0012793465,0.00041992805,0.0011806199,0.00067223207,0.00063941564,0.0022381244],"category_scores_gemma":[0.0053016134,0.0002764002,0.00060616236,0.00084156025,0.00030502482,0.0019859937,0.0009955444,0.001079397,0.0016458365],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006912819,0.00019070666,0.001089898,0.0005790553,0.000111362606,0.00018376498,0.00080631045,0.019903721,0.099534936,0.0021618234,0.007978269,0.86676884],"study_design_scores_gemma":[0.00024237881,0.0011263989,0.0048763356,0.000096665746,0.0004927271,0.00028989394,0.0008457529,0.77570814,0.18230326,0.007564406,0.026342466,0.00011157887],"about_ca_topic_score_codex":0.0013615212,"about_ca_topic_score_gemma":0.002223882,"teacher_disagreement_score":0.0022381244,"about_ca_system_score_codex":0.00029953866,"about_ca_system_score_gemma":0.00048419647,"threshold_uncertainty_score":0.007487297},"labels":[],"label_agreement":null},{"id":"W4393159548","doi":"10.1609/aaai.v38i12.29263","title":"Narrowing the Gap between Supervised and Unsupervised Sentence Representation Learning with Large Language Model","year":2024,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Fundamental Research Funds for the Central Universities; State Key Laboratory of Software Development Environment; National Natural Science Foundation of China","keywords":"Sentence; Natural language processing; Artificial intelligence; Computer science; Representation (politics); Unsupervised learning; Language model; Linguistics; Machine learning; Philosophy","score_opus":0.09381631991397431,"score_gpt":0.3127379495608163,"score_spread":0.21892162964684198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393159548","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36878905,0.002009144,0.61422354,0.0021328456,0.00017026407,0.00028976682,0.0006650152,0.0064122537,0.0053081857],"genre_scores_gemma":[0.818263,0.00035696683,0.17584325,0.00063190033,0.00009881634,0.0003080576,0.0023664602,0.00065436185,0.0014771825],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9898784,0.0057236706,0.00051204045,0.0022812209,0.0012826876,0.00032199608],"domain_scores_gemma":[0.94565076,0.04243835,0.0014663441,0.0076956996,0.0019650804,0.0007836594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010861003,0.0016224857,0.0016411024,0.0011689807,0.0009732795,0.0019926263,0.002250228,0.002049373,0.0015252278],"category_scores_gemma":[0.05164589,0.0007202334,0.0009083318,0.0009391843,0.0013836187,0.00812452,0.004919515,0.004367847,0.0011283342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015232433,0.0019575441,0.013238631,0.000758388,0.00039255546,0.00035056437,0.00200546,0.21811877,0.022901356,0.014837943,0.010718808,0.7131967],"study_design_scores_gemma":[0.000059188595,0.00049262017,0.002111769,0.000035323916,0.000032204964,0.00012459458,0.0002451566,0.96514124,0.010301918,0.019785237,0.0016192121,0.000051503204],"about_ca_topic_score_codex":0.0018166034,"about_ca_topic_score_gemma":0.0034023605,"teacher_disagreement_score":0.010861003,"about_ca_system_score_codex":0.0010493296,"about_ca_system_score_gemma":0.0014162206,"threshold_uncertainty_score":0.057439208},"labels":[],"label_agreement":null},{"id":"W4393160884","doi":"10.22541/essoar.171136837.71755629/v1","title":"Potential Benefits and Dangers of Using Large Language Models for Advancing Sustainability Science and Communication","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Sustainability; Unintended consequences; Sustainable development; Sustainability science; Engineering ethics; Political science; Social sustainability; Engineering; Ecology","score_opus":0.0228580133462147,"score_gpt":0.30575226402843503,"score_spread":0.28289425068222035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393160884","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0139347715,0.0044388794,0.8947317,0.062460184,0.0008124333,0.00023434624,0.0007079344,0.0014633014,0.021216502],"genre_scores_gemma":[0.31195912,0.0061605414,0.6630358,0.0063722003,0.0015904868,0.0011639408,0.0015517767,0.0012761245,0.0068900175],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9147944,0.07141584,0.0027055622,0.0034181082,0.0067021,0.0009639016],"domain_scores_gemma":[0.66614246,0.2893198,0.006213489,0.02708058,0.00893835,0.0023054024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07205562,0.0015174927,0.0016020127,0.003421182,0.00274978,0.016505241,0.0041508614,0.00453757,0.0062404484],"category_scores_gemma":[0.19150658,0.0016374375,0.0024791392,0.003443081,0.008958357,0.03608446,0.011284355,0.009334131,0.0029158732],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012755267,0.000067911285,0.0022438527,0.00048843643,0.00013616125,0.00017647196,0.0040815533,0.020561693,0.0006122371,0.90995556,0.007835222,0.0537133],"study_design_scores_gemma":[0.000031924643,0.000030497184,0.00022357362,0.00031005713,0.00004827668,0.000119458586,0.0009428532,0.068879366,0.0006204346,0.88905764,0.039662957,0.00007295041],"about_ca_topic_score_codex":0.0070743146,"about_ca_topic_score_gemma":0.009456172,"teacher_disagreement_score":0.07205562,"about_ca_system_score_codex":0.005629371,"about_ca_system_score_gemma":0.008141407,"threshold_uncertainty_score":0.3810711},"labels":[],"label_agreement":null},{"id":"W4393183687","doi":"10.1145/3639279","title":"DTT: An Example-Driven Tabular Transformer for Joinability by Leveraging Large Language Models","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on Management of Data","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Transformer; Computer science; Engineering; Electrical engineering; Voltage","score_opus":0.10344835003044657,"score_gpt":0.3135535429041299,"score_spread":0.21010519287368334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393183687","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014838665,0.00029568336,0.9390623,0.00038921632,0.00015148318,0.00017716984,0.0032145034,0.038470205,0.003400696],"genre_scores_gemma":[0.2596592,0.00047294633,0.70783186,0.00063300977,0.00012681946,0.00034707744,0.018027607,0.003564385,0.009337069],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990103,0.00018874303,0.00008800172,0.0002983102,0.0003332073,0.00008137789],"domain_scores_gemma":[0.9977513,0.0008313867,0.00015925051,0.0008287902,0.0003334803,0.000095785836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001350108,0.0010169699,0.00073226116,0.0014301519,0.0005445267,0.0017739859,0.002506195,0.0009190601,0.01073176],"category_scores_gemma":[0.006887335,0.00054995547,0.0017535323,0.0016715825,0.000907535,0.006842368,0.002464308,0.0024246697,0.005227981],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005623537,0.0004250442,0.0043151462,0.0004880039,0.00012025084,0.00039654126,0.00050036714,0.10527438,0.013235885,0.056524213,0.066849574,0.7513082],"study_design_scores_gemma":[0.000054082855,0.00008074307,0.00027980554,0.00003499546,0.000032562104,0.00016119111,0.00009108228,0.9183519,0.0131200515,0.048556384,0.01920556,0.000031634372],"about_ca_topic_score_codex":0.0064290254,"about_ca_topic_score_gemma":0.0109417755,"teacher_disagreement_score":0.01073176,"about_ca_system_score_codex":0.0011175158,"about_ca_system_score_gemma":0.0022653006,"threshold_uncertainty_score":0.035901308},"labels":[],"label_agreement":null},{"id":"W4393212490","doi":"10.48550/arxiv.2403.15567","title":"Do not trust what you trust: Miscalibration in Semi-supervised Learning","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Psychology; Knowledge management; Computer science; Artificial intelligence","score_opus":0.06709897106881242,"score_gpt":0.19518362559078378,"score_spread":0.12808465452197138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393212490","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08872207,0.0017040519,0.90102583,0.0030300352,0.00015902842,0.00014931703,0.00028204205,0.001800495,0.003127061],"genre_scores_gemma":[0.8669119,0.000390061,0.12830582,0.0012210128,0.00029322002,0.00020044927,0.0005038521,0.00036902243,0.0018046551],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9646642,0.021049283,0.001682298,0.0062106936,0.0055711702,0.0008224416],"domain_scores_gemma":[0.8370919,0.11296851,0.014163667,0.025317991,0.00878136,0.0016765596],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030482225,0.0016384466,0.0021579692,0.0020715925,0.0023025442,0.0038390392,0.004502131,0.0038043875,0.0012443151],"category_scores_gemma":[0.11821052,0.0012977646,0.00088928104,0.0021728023,0.006522851,0.007801736,0.0053433776,0.0058868984,0.000732457],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021105015,0.0003684622,0.035464745,0.0012762259,0.0009608291,0.0010805376,0.0063716252,0.3997951,0.008601395,0.08842072,0.01572426,0.43982568],"study_design_scores_gemma":[0.00006141674,0.00014286574,0.0027842503,0.00017072899,0.00007650765,0.00038839862,0.0003431655,0.86565167,0.007537099,0.11838705,0.0043677767,0.00008913042],"about_ca_topic_score_codex":0.0037939728,"about_ca_topic_score_gemma":0.0045484668,"teacher_disagreement_score":0.030482225,"about_ca_system_score_codex":0.0025738508,"about_ca_system_score_gemma":0.002281309,"threshold_uncertainty_score":0.16120732},"labels":[],"label_agreement":null},{"id":"W4393319378","doi":"10.21203/rs.3.rs-4169544/v1","title":"Integrating Quantum CI with Generative AI for Taiwanese/English Co-Learning: TAIDE-based Knowledge Graph Construction and Multimodal Data Transformation","year":2024,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Science and Technology Council","keywords":"Generative grammar; Transformation (genetics); Graph; Computer science; Artificial intelligence; Natural language processing; Theoretical computer science","score_opus":0.08285924914743456,"score_gpt":0.3999924440936812,"score_spread":0.3171331949462466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393319378","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047724884,0.00028948847,0.94069517,0.0006951103,0.000058704492,0.00011818779,0.00035633624,0.0016822699,0.008379691],"genre_scores_gemma":[0.7226864,0.00020765097,0.27178618,0.00016056035,0.000036912115,0.00016723218,0.0008642228,0.0002809957,0.003809893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99919254,0.0003454869,0.000049122973,0.00022999829,0.00011084009,0.000071964176],"domain_scores_gemma":[0.99703515,0.0017424871,0.0000929165,0.000637768,0.00035593988,0.00013571176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013557258,0.00038616205,0.0005754256,0.001457077,0.00080991327,0.0016423889,0.0012811256,0.0006253926,0.004941051],"category_scores_gemma":[0.006257935,0.00032398984,0.0009782722,0.0018384036,0.0010212273,0.0033024584,0.0030504875,0.0016580572,0.0008637762],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026483298,0.00039389133,0.0061577046,0.00039735483,0.00024307077,0.00035760214,0.0024991417,0.08284217,0.018919788,0.17546114,0.006892459,0.7055708],"study_design_scores_gemma":[0.000018754921,0.0000472516,0.0017027057,0.00003480757,0.000056077788,0.000078245066,0.00060004275,0.8390718,0.0066967285,0.146783,0.004879337,0.00003133355],"about_ca_topic_score_codex":0.010496101,"about_ca_topic_score_gemma":0.015659614,"teacher_disagreement_score":0.010496101,"about_ca_system_score_codex":0.0011940543,"about_ca_system_score_gemma":0.0018169726,"threshold_uncertainty_score":0.02087003},"labels":[],"label_agreement":null},{"id":"W4393335480","doi":"10.1038/s41746-024-01074-z","title":"Foundation metrics for evaluating effectiveness of healthcare conversations powered by generative AI","year":2024,"lang":"en","type":"article","venue":"npj Digital Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":184,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institute of Standards and Technology","keywords":"Health care; Computer science; Personalization; Set (abstract data type); Process (computing); Comprehension; Human–computer interaction; World Wide Web","score_opus":0.059876822644538336,"score_gpt":0.3823261757586979,"score_spread":0.32244935311415956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393335480","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41735488,0.008741854,0.5330786,0.0014338134,0.0005247866,0.004172828,0.005899128,0.004469496,0.024324555],"genre_scores_gemma":[0.78426945,0.00082971575,0.20594175,0.0002527768,0.00011977512,0.0033127582,0.0036794369,0.0003346324,0.0012597373],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95976686,0.022220816,0.0052103787,0.0025988815,0.009133103,0.0010699857],"domain_scores_gemma":[0.81317097,0.14794819,0.011686477,0.009905272,0.014932795,0.0023562538],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025364084,0.0023678297,0.0013611115,0.006108721,0.0010805009,0.003921134,0.0019083532,0.002132912,0.0026644773],"category_scores_gemma":[0.1398605,0.0004638364,0.0013504653,0.0030511545,0.0016286257,0.0033381998,0.002929462,0.0016574046,0.0009015829],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035259034,0.0018940445,0.08912361,0.005973418,0.0014687937,0.00029116528,0.004010797,0.15956347,0.0285413,0.018690486,0.008830114,0.6780869],"study_design_scores_gemma":[0.0002811993,0.008378941,0.08310831,0.0013644253,0.000888885,0.0008260967,0.002952256,0.7992863,0.053500995,0.03228115,0.016624972,0.0005064644],"about_ca_topic_score_codex":0.002596543,"about_ca_topic_score_gemma":0.001978738,"teacher_disagreement_score":0.025364084,"about_ca_system_score_codex":0.0019690138,"about_ca_system_score_gemma":0.0016655283,"threshold_uncertainty_score":0.13413972},"labels":[],"label_agreement":null},{"id":"W4393493608","doi":"10.5281/zenodo.3930995","title":"Preprocessing scripts and data for study: Identifying high-confidence capture Hi-C interactions using CHiCANE","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Cancer Research","funders":"","keywords":"Scripting language; Preprocessor; Computer science; Artificial intelligence; Programming language","score_opus":0.160977707385079,"score_gpt":0.3419859801446908,"score_spread":0.1810082727596118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393493608","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013324919,0.0001243491,0.001509354,0.000052070027,0.00003161594,0.00010238691,0.9916889,0.0044368496,0.0007221377],"genre_scores_gemma":[0.0012248089,0.000040492283,0.0025308381,0.000044026747,0.0000062221893,0.00040024298,0.995,0.00025676354,0.00049654365],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990841,0.00012877153,0.00010547298,0.00041068066,0.00015628309,0.00011468183],"domain_scores_gemma":[0.9983041,0.0006045254,0.00013137131,0.00046995838,0.00033658138,0.00015342906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013449205,0.0028929545,0.0012359045,0.0027889062,0.00091513444,0.0016187254,0.0020758673,0.0015635503,0.030544043],"category_scores_gemma":[0.0040552123,0.0006469366,0.0013576839,0.0032599308,0.00042180644,0.00061302475,0.0010827014,0.0014829428,0.038359087],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042147873,0.00013315518,0.004335882,0.0017381398,0.000106169166,0.00011453937,0.00008391051,0.0013624079,0.0041001816,0.00079533947,0.97573984,0.0110690445],"study_design_scores_gemma":[0.00083478924,0.00018238873,0.01919672,0.00031363647,0.00023019032,0.00044004526,0.00022023603,0.0081079025,0.012353952,0.004623439,0.9533881,0.000108563036],"about_ca_topic_score_codex":0.011743771,"about_ca_topic_score_gemma":0.023113752,"teacher_disagreement_score":0.030544043,"about_ca_system_score_codex":0.0010710066,"about_ca_system_score_gemma":0.0023218864,"threshold_uncertainty_score":0.102180004},"labels":[],"label_agreement":null},{"id":"W4393626689","doi":"10.5281/zenodo.3381673","title":"Cross-Domain Modeling of Sentence-Level Evidence for Document Retrieval","year":2019,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Information retrieval; Sentence; Computer science; Domain (mathematical analysis); Natural language processing; Artificial intelligence; Mathematics","score_opus":0.2650748796003983,"score_gpt":0.382171404372805,"score_spread":0.11709652477240667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393626689","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030917946,0.0028832871,0.037070893,0.0019765724,0.0011674262,0.0006034305,0.8907994,0.024865393,0.0097156465],"genre_scores_gemma":[0.03245374,0.0004402584,0.017137403,0.0002748608,0.00010983276,0.00044313836,0.9429979,0.0007004814,0.0054424205],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99874735,0.00046011995,0.000094563446,0.00040006233,0.00020033386,0.00009757703],"domain_scores_gemma":[0.99527466,0.0023322708,0.0001505871,0.0013490685,0.00070155616,0.00019186705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034704714,0.0025379937,0.0012303835,0.0023722234,0.00087676715,0.0019186513,0.0029355937,0.002227176,0.034338854],"category_scores_gemma":[0.0141953025,0.00078112894,0.0029389048,0.002208271,0.00054309994,0.0023877895,0.0016922558,0.0036670694,0.045505725],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010222236,0.00042942964,0.0050068456,0.0013221985,0.00039331533,0.00015504641,0.0000615757,0.021319957,0.0017187075,0.0017457815,0.9056665,0.061158426],"study_design_scores_gemma":[0.0020871821,0.0012107454,0.032771617,0.0009455301,0.0010341562,0.0012753693,0.00036685631,0.50113755,0.019839443,0.032488987,0.40650228,0.00034034633],"about_ca_topic_score_codex":0.011743582,"about_ca_topic_score_gemma":0.019330407,"teacher_disagreement_score":0.034338854,"about_ca_system_score_codex":0.0015445278,"about_ca_system_score_gemma":0.0015202087,"threshold_uncertainty_score":0.11487496},"labels":[],"label_agreement":null},{"id":"W4393638019","doi":"10.5281/zenodo.3630694","title":"Russian News 2017","year":2020,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Optech (Canada)","funders":"","keywords":"Information retrieval; Computer science","score_opus":0.09046711285540973,"score_gpt":0.2858571688788321,"score_spread":0.19539005602342238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393638019","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014515143,0.00023677712,0.0004124771,0.00014498431,0.00017026904,0.000050951196,0.9929482,0.0016192304,0.0029656168],"genre_scores_gemma":[0.0009428918,0.00005045865,0.0005574642,0.000043224856,0.000018469587,0.00008759276,0.99663866,0.00010371587,0.0015574757],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983487,0.00042371644,0.00014618077,0.00043732068,0.00037717147,0.00026698445],"domain_scores_gemma":[0.99740654,0.0006756565,0.0001502713,0.0006672743,0.0007255583,0.00037469843],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014405607,0.0031881023,0.0014157783,0.004815359,0.0015453082,0.0020682083,0.0020204692,0.0024459136,0.047499795],"category_scores_gemma":[0.0055865278,0.00064423337,0.0017202647,0.004204505,0.00053840445,0.0016637762,0.003122017,0.0018902996,0.095664375],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013280379,0.000052539723,0.0006507077,0.0004338891,0.00003732545,0.000056120003,0.00007530749,0.00029307493,0.00038824455,0.0003617705,0.9922672,0.005251027],"study_design_scores_gemma":[0.00023627045,0.000054960437,0.0071804635,0.000222307,0.00008364448,0.00019555983,0.00037559538,0.0018526017,0.0013465367,0.0011804694,0.98720807,0.00006351011],"about_ca_topic_score_codex":0.02175567,"about_ca_topic_score_gemma":0.046336308,"teacher_disagreement_score":0.047499795,"about_ca_system_score_codex":0.0012137585,"about_ca_system_score_gemma":0.0023669403,"threshold_uncertainty_score":0.1589027},"labels":[],"label_agreement":null},{"id":"W4393650160","doi":"10.5281/zenodo.10446176","title":"On Inter-dataset Code Duplication and Data Leakage in Large Language Models","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; McGill University","funders":"","keywords":"Computer science; Code (set theory); Leakage (economics); Programming language","score_opus":0.07910873542123568,"score_gpt":0.3116455489015773,"score_spread":0.2325368134803416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393650160","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008795487,0.00072332943,0.005848454,0.00084268587,0.00022490798,0.000096653894,0.96811664,0.011171767,0.004180044],"genre_scores_gemma":[0.006304898,0.00016483966,0.0035691569,0.00013256448,0.000022308092,0.00011314407,0.988143,0.000482653,0.0010672768],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.99691725,0.00077902945,0.00028829864,0.0008760104,0.00090197567,0.00023743993],"domain_scores_gemma":[0.99236226,0.0032772413,0.0004245881,0.002531594,0.0011186711,0.00028572246],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0023449222,0.0019901244,0.0009698951,0.00390733,0.0014239737,0.001961066,0.002444668,0.0016373043,0.009046121],"category_scores_gemma":[0.014476722,0.00045609436,0.0017356673,0.0060979347,0.00073565694,0.0023454085,0.0023705089,0.00204448,0.0135899205],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001424141,0.00007927966,0.0030331719,0.0009089563,0.000098672106,0.0001337039,0.000104536455,0.0036968605,0.0006924811,0.0029288589,0.9709216,0.017259516],"study_design_scores_gemma":[0.0003971545,0.00012565672,0.013660354,0.00048696733,0.00013791375,0.00074617926,0.00033416424,0.0335048,0.0054578525,0.01813847,0.9268971,0.000113468785],"about_ca_topic_score_codex":0.017484153,"about_ca_topic_score_gemma":0.044307232,"teacher_disagreement_score":0.9976551,"about_ca_system_score_codex":0.0017857739,"about_ca_system_score_gemma":0.0026437505,"threshold_uncertainty_score":0.034764767},"labels":[],"label_agreement":null},{"id":"W4393704820","doi":"10.5281/zenodo.3372763","title":"Cross-Domain Modeling of Sentence-Level Evidence for Document Retrieval","year":2019,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Information retrieval; Sentence; Computer science; Domain (mathematical analysis); Natural language processing; Artificial intelligence; Mathematics","score_opus":0.1617745458923477,"score_gpt":0.33566998973817835,"score_spread":0.17389544384583067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393704820","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03279457,0.0029785745,0.027415551,0.0014977021,0.0011519307,0.0006548772,0.9059013,0.01994808,0.007657338],"genre_scores_gemma":[0.02255678,0.00034440178,0.011544721,0.00019459536,0.00008966616,0.00036823968,0.9603424,0.00040268703,0.004156603],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99825734,0.00062667177,0.00014086331,0.000525608,0.00030732623,0.00014205372],"domain_scores_gemma":[0.99568766,0.0018982036,0.00018169054,0.0012946833,0.00072070496,0.00021699874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036945676,0.0030222537,0.0015215063,0.0022998883,0.00092652213,0.0018588708,0.003156278,0.0023490498,0.022385694],"category_scores_gemma":[0.011711473,0.00080132275,0.0028135667,0.002113356,0.0006155138,0.0019847758,0.0017744167,0.0036626556,0.037139717],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010918239,0.0005738674,0.0041861893,0.0015617769,0.00039537626,0.00019900153,0.0000617225,0.016719544,0.0023246442,0.0012994312,0.9122641,0.059322394],"study_design_scores_gemma":[0.0024888294,0.0016768858,0.034788173,0.00097616285,0.0011669076,0.0018778772,0.0004323743,0.43617922,0.027631635,0.024922676,0.4674665,0.00039275954],"about_ca_topic_score_codex":0.011976289,"about_ca_topic_score_gemma":0.019261902,"teacher_disagreement_score":0.022385694,"about_ca_system_score_codex":0.0016495577,"about_ca_system_score_gemma":0.0016338814,"threshold_uncertainty_score":0.07488757},"labels":[],"label_agreement":null},{"id":"W4393716831","doi":"10.1145/3654660","title":"Computational Politeness in Natural Language Processing: A Survey","year":2024,"lang":"en","type":"review","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Ministry of Electronics and Information technology","keywords":"Computer science; Politeness; Natural language processing; Natural language; Natural (archaeology); Artificial intelligence; Linguistics","score_opus":0.08069860879599831,"score_gpt":0.38061836495845297,"score_spread":0.29991975616245464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393716831","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026091516,0.9371058,0.039266784,0.006133159,0.000542934,0.00012059683,0.0002644959,0.00042947973,0.013527629],"genre_scores_gemma":[0.020393647,0.9498028,0.023534972,0.0011607056,0.0018638179,0.00018925955,0.00092591817,0.00014559814,0.0019832933],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978594,0.00072769023,0.00019666983,0.0004392279,0.0006654467,0.00011160528],"domain_scores_gemma":[0.98149925,0.016099863,0.00031778982,0.0007153624,0.001190891,0.0001768642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004412966,0.0011548987,0.0017387717,0.0050561517,0.00069435645,0.004179808,0.0017123619,0.0017718138,0.0051307855],"category_scores_gemma":[0.014241444,0.0008039131,0.001246999,0.00788621,0.0018583047,0.008606941,0.0019443533,0.0026962622,0.0030558868],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000054621694,0.000115766874,0.0012910351,0.007861344,0.000089375666,0.000044896828,0.00038865348,0.0018335278,0.00067297864,0.03959123,0.029629793,0.91842675],"study_design_scores_gemma":[0.000038278307,0.00017498691,0.005556418,0.0070216805,0.00017803011,0.000775277,0.0009761988,0.015478968,0.001688323,0.12028112,0.8477203,0.000110456465],"about_ca_topic_score_codex":0.0024819237,"about_ca_topic_score_gemma":0.0019817452,"teacher_disagreement_score":0.0051307855,"about_ca_system_score_codex":0.0017218326,"about_ca_system_score_gemma":0.002780482,"threshold_uncertainty_score":0.023338258},"labels":[],"label_agreement":null},{"id":"W4393717402","doi":"10.5281/zenodo.3372764","title":"Cross-Domain Modeling of Sentence-Level Evidence for Document Retrieval","year":2019,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Information retrieval; Sentence; Computer science; Domain (mathematical analysis); Natural language processing; Artificial intelligence; Mathematics","score_opus":0.2650748796003983,"score_gpt":0.382171404372805,"score_spread":0.11709652477240667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393717402","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030917946,0.0028832871,0.037070893,0.0019765724,0.0011674262,0.0006034305,0.8907994,0.024865393,0.0097156465],"genre_scores_gemma":[0.03245374,0.0004402584,0.017137403,0.0002748608,0.00010983276,0.00044313836,0.9429979,0.0007004814,0.0054424205],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99874735,0.00046011995,0.000094563446,0.00040006233,0.00020033386,0.00009757703],"domain_scores_gemma":[0.99527466,0.0023322708,0.0001505871,0.0013490685,0.00070155616,0.00019186705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034704714,0.0025379937,0.0012303835,0.0023722234,0.00087676715,0.0019186513,0.0029355937,0.002227176,0.034338854],"category_scores_gemma":[0.0141953025,0.00078112894,0.0029389048,0.002208271,0.00054309994,0.0023877895,0.0016922558,0.0036670694,0.045505725],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010222236,0.00042942964,0.0050068456,0.0013221985,0.00039331533,0.00015504641,0.0000615757,0.021319957,0.0017187075,0.0017457815,0.9056665,0.061158426],"study_design_scores_gemma":[0.0020871821,0.0012107454,0.032771617,0.0009455301,0.0010341562,0.0012753693,0.00036685631,0.50113755,0.019839443,0.032488987,0.40650228,0.00034034633],"about_ca_topic_score_codex":0.011743582,"about_ca_topic_score_gemma":0.019330407,"teacher_disagreement_score":0.034338854,"about_ca_system_score_codex":0.0015445278,"about_ca_system_score_gemma":0.0015202087,"threshold_uncertainty_score":0.11487496},"labels":[],"label_agreement":null},{"id":"W4393773180","doi":"10.5281/zenodo.6149598","title":"TopiOCQA processed Wikipedia data","year":2022,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Information retrieval; Data science; World Wide Web; Data mining","score_opus":0.08125994892770544,"score_gpt":0.2764849832977754,"score_spread":0.19522503437006994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393773180","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014834012,0.000107853615,0.0003949667,0.00007265762,0.000075478514,0.000099455625,0.99388146,0.0025367935,0.0013478644],"genre_scores_gemma":[0.0009802403,0.000027353199,0.0008961455,0.000023118297,0.000007435522,0.00011624083,0.9971418,0.00017433117,0.00063333235],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976579,0.00035012516,0.0002467108,0.0006678156,0.00073496473,0.00034245546],"domain_scores_gemma":[0.99623823,0.00062359456,0.00021086457,0.0009392629,0.0015967758,0.000391222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015968918,0.0034006194,0.0013822727,0.007145502,0.0013256498,0.002466022,0.0025559282,0.0021691653,0.026307888],"category_scores_gemma":[0.007030532,0.0008264021,0.0022865215,0.006829483,0.00074337656,0.0013553145,0.0019308737,0.0023847746,0.046761766],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019228455,0.00016990185,0.0013524891,0.00088035635,0.0001044591,0.00007312993,0.00008017348,0.001273487,0.0015745515,0.0006855915,0.98721236,0.0064012664],"study_design_scores_gemma":[0.00059860037,0.00010385026,0.008313859,0.00033545948,0.00015745989,0.00016611598,0.00030268688,0.006448708,0.0057375883,0.0019755736,0.9757357,0.00012440539],"about_ca_topic_score_codex":0.050622195,"about_ca_topic_score_gemma":0.06532497,"teacher_disagreement_score":0.050622195,"about_ca_system_score_codex":0.001787672,"about_ca_system_score_gemma":0.004830393,"threshold_uncertainty_score":0.10065508},"labels":[],"label_agreement":null},{"id":"W4393823895","doi":"10.5281/zenodo.3932441","title":"SemEval-2020 Task 5: Modelling Causal Reasoning in Language: Detecting Counterfactuals","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; Queen's University","funders":"","keywords":"Counterfactual conditional; Computer science; Task (project management); SemEval; Causal reasoning; Natural language processing; Artificial intelligence; Linguistics; Psychology; Counterfactual thinking; Philosophy; Cognition; Engineering","score_opus":0.03623281979995987,"score_gpt":0.25938332611982257,"score_spread":0.2231505063198627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393823895","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01695245,0.0019316542,0.008746411,0.0014420372,0.0005225582,0.00073778105,0.93594265,0.022780068,0.010944353],"genre_scores_gemma":[0.010551383,0.00015962725,0.010093204,0.00026485196,0.000039257142,0.0005298769,0.9756564,0.00035973825,0.0023457166],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974891,0.0008363956,0.00022055335,0.00075132767,0.00048322778,0.00021937881],"domain_scores_gemma":[0.99591184,0.002143738,0.00022571142,0.000836653,0.00056268927,0.0003193519],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033384184,0.0062356275,0.0015470728,0.0027277507,0.0016326902,0.0033755903,0.0058928393,0.005091087,0.030320646],"category_scores_gemma":[0.009591959,0.0010169282,0.0037480635,0.0021246893,0.0010722296,0.003409224,0.0032527898,0.0037034964,0.033991627],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006743952,0.00042968424,0.002525527,0.0024309645,0.00021775306,0.00034718114,0.00017782576,0.005114592,0.0014655627,0.0026101128,0.95672995,0.027276576],"study_design_scores_gemma":[0.0034632522,0.00061978307,0.01629503,0.0013073711,0.00045528106,0.0016981658,0.0011824125,0.111200325,0.012628787,0.026257655,0.8246062,0.0002857005],"about_ca_topic_score_codex":0.027014196,"about_ca_topic_score_gemma":0.06076585,"teacher_disagreement_score":0.030320646,"about_ca_system_score_codex":0.0033782802,"about_ca_system_score_gemma":0.003299796,"threshold_uncertainty_score":0.10143268},"labels":[],"label_agreement":null},{"id":"W4393924586","doi":"10.48550/arxiv.2404.01399","title":"Developing Safe and Responsible Large Language Model : Can We Balance Bias Reduction and Language Understanding in Large Language Models?","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Government of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Business","score_opus":0.12962998173143916,"score_gpt":0.23450495749436628,"score_spread":0.10487497576292712,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393924586","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07686141,0.00207396,0.87803817,0.0030451748,0.0003124571,0.00032162765,0.0030672064,0.032198712,0.0040813233],"genre_scores_gemma":[0.4420123,0.0008770543,0.53156483,0.0026783464,0.00018836767,0.0009735117,0.012781572,0.0041560926,0.004767777],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99720436,0.001433384,0.0001559074,0.0007378263,0.0002979445,0.00017052676],"domain_scores_gemma":[0.99016225,0.005683968,0.00035944238,0.0026053942,0.00090956443,0.00027933586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005762363,0.0020609323,0.0013236254,0.0009661885,0.00068371213,0.0025813668,0.0033866,0.0020733096,0.003588965],"category_scores_gemma":[0.032818835,0.0008852244,0.0015693471,0.0009315038,0.0015158751,0.0072335238,0.0030424197,0.0049544037,0.0049361675],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096313236,0.0005018561,0.013464168,0.0011020937,0.00052071136,0.00047676393,0.001411578,0.2724908,0.025004921,0.025050206,0.041285783,0.617728],"study_design_scores_gemma":[0.0001786942,0.00015098807,0.000737806,0.00010233106,0.00010571391,0.00013550239,0.00022591671,0.9215601,0.015311238,0.050603088,0.010821079,0.0000675484],"about_ca_topic_score_codex":0.0068962746,"about_ca_topic_score_gemma":0.016085869,"teacher_disagreement_score":0.0068962746,"about_ca_system_score_codex":0.0014203084,"about_ca_system_score_gemma":0.0032216327,"threshold_uncertainty_score":0.030474663},"labels":[],"label_agreement":null},{"id":"W4394006050","doi":"10.11591/ijeecs.v34.i3.pp1760-1769","title":"Exploring the intricacies of human memory and its analogous representation in ChatGPT","year":2024,"lang":"en","type":"article","venue":"Indonesian Journal of Electrical Engineering and Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Memorization; Recall; Computer science; Parallels; Human memory; Relevance (law); Representation (politics); Cognitive psychology; Mental representation; Cognitive science; Human–computer interaction; Psychology; Cognition; Neuroscience","score_opus":0.03377260314177536,"score_gpt":0.25043493885588186,"score_spread":0.2166623357141065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394006050","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6583748,0.0011516007,0.2979942,0.0011728641,0.000049070917,0.00009014238,0.00028082656,0.0006787003,0.04020785],"genre_scores_gemma":[0.9772694,0.00019219519,0.02026408,0.000037946815,0.000008761474,0.000033595807,0.00007558913,0.000028429171,0.002089991],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99977475,0.00008135515,0.000010417968,0.000071865965,0.0000352135,0.000026501648],"domain_scores_gemma":[0.99860555,0.00081103004,0.000109974055,0.00028492988,0.00012449197,0.00006393879],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00044488304,0.0002787428,0.0002620034,0.00062376034,0.0003754053,0.002571016,0.00079775095,0.00058969925,0.0031375817],"category_scores_gemma":[0.0034818735,0.00027588423,0.00048389588,0.00063019653,0.0018783128,0.004850698,0.0010082005,0.0005350395,0.00030608778],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081530743,0.00021730927,0.038102858,0.0009139483,0.0001800968,0.0033641588,0.033466935,0.086741604,0.059451085,0.56258905,0.002549716,0.21160802],"study_design_scores_gemma":[0.0000656718,0.00032291486,0.026781354,0.00019898439,0.0001527518,0.0026406152,0.0068794456,0.50911963,0.019554904,0.41588414,0.018283878,0.000115732546],"about_ca_topic_score_codex":0.003981183,"about_ca_topic_score_gemma":0.0022539631,"teacher_disagreement_score":0.003981183,"about_ca_system_score_codex":0.0007580991,"about_ca_system_score_gemma":0.00045138205,"threshold_uncertainty_score":0.010496259},"labels":[],"label_agreement":null},{"id":"W4394039396","doi":"10.5281/zenodo.3930996","title":"Preprocessing scripts and data for study: Identifying high-confidence capture Hi-C interactions using CHiCANE","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Cancer Research","funders":"","keywords":"Preprocessor; Scripting language; Computer science; Data pre-processing; Artificial intelligence; Operating system","score_opus":0.160977707385079,"score_gpt":0.3419859801446908,"score_spread":0.1810082727596118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394039396","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013324919,0.0001243491,0.001509354,0.000052070027,0.00003161594,0.00010238691,0.9916889,0.0044368496,0.0007221377],"genre_scores_gemma":[0.0012248089,0.000040492283,0.0025308381,0.000044026747,0.0000062221893,0.00040024298,0.995,0.00025676354,0.00049654365],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990841,0.00012877153,0.00010547298,0.00041068066,0.00015628309,0.00011468183],"domain_scores_gemma":[0.9983041,0.0006045254,0.00013137131,0.00046995838,0.00033658138,0.00015342906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013449205,0.0028929545,0.0012359045,0.0027889062,0.00091513444,0.0016187254,0.0020758673,0.0015635503,0.030544043],"category_scores_gemma":[0.0040552123,0.0006469366,0.0013576839,0.0032599308,0.00042180644,0.00061302475,0.0010827014,0.0014829428,0.038359087],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042147873,0.00013315518,0.004335882,0.0017381398,0.000106169166,0.00011453937,0.00008391051,0.0013624079,0.0041001816,0.00079533947,0.97573984,0.0110690445],"study_design_scores_gemma":[0.00083478924,0.00018238873,0.01919672,0.00031363647,0.00023019032,0.00044004526,0.00022023603,0.0081079025,0.012353952,0.004623439,0.9533881,0.000108563036],"about_ca_topic_score_codex":0.011743771,"about_ca_topic_score_gemma":0.023113752,"teacher_disagreement_score":0.030544043,"about_ca_system_score_codex":0.0010710066,"about_ca_system_score_gemma":0.0023218864,"threshold_uncertainty_score":0.102180004},"labels":[],"label_agreement":null},{"id":"W4394169281","doi":"10.6084/m9.figshare.24955719","title":"Identification and Description of Emotions by Current Large Language Models - Dataset","year":2024,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Identification (biology); Current (fluid); Computer science; Natural language processing; Psychology; Cognitive psychology; Cognitive science; Geology; Biology; Oceanography; Ecology","score_opus":0.07209381814899139,"score_gpt":0.3118934366214614,"score_spread":0.23979961847247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394169281","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025529247,0.0017597977,0.005932182,0.0009334671,0.0002577013,0.00030465864,0.95355594,0.0063423146,0.0053846575],"genre_scores_gemma":[0.014842947,0.00019046682,0.0055027143,0.00010983959,0.000023581804,0.00031604024,0.9774799,0.0001184421,0.0014160422],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99811965,0.0006696448,0.0002436806,0.00051133917,0.00032108987,0.00013454627],"domain_scores_gemma":[0.9974482,0.0012752445,0.00015368588,0.00056603126,0.0004095828,0.00014721278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022783612,0.0023107633,0.0006753606,0.002975795,0.0007957504,0.0016026191,0.002306819,0.0020860557,0.0078335805],"category_scores_gemma":[0.007828608,0.00033099513,0.0018073071,0.0027150991,0.0005167713,0.0014399525,0.0015996059,0.0018502902,0.014885029],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050852355,0.00033980224,0.013813733,0.0019522528,0.00026904038,0.00046009017,0.0004047141,0.006099342,0.0025163025,0.0020337682,0.9178579,0.05374456],"study_design_scores_gemma":[0.0006678707,0.00029695017,0.044285182,0.00051987806,0.00029465472,0.0015567819,0.0011645433,0.05933388,0.0077001574,0.0063176304,0.87765515,0.00020720305],"about_ca_topic_score_codex":0.015523194,"about_ca_topic_score_gemma":0.031854283,"teacher_disagreement_score":0.015523194,"about_ca_system_score_codex":0.0015567614,"about_ca_system_score_gemma":0.001476561,"threshold_uncertainty_score":0.03086567},"labels":[],"label_agreement":null},{"id":"W4394208231","doi":"10.6084/m9.figshare.13368924.v1","title":"Topic Modeling on Triage Notes with Semi-orthogonal Non-negative Matrix Factorization","year":2020,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Triage; Factorization; Computer science; Matrix decomposition; Matrix (chemical analysis); Combinatorics; Mathematics; Medicine; Algorithm; Medical emergency; Physics; Chemistry","score_opus":0.059388701860119894,"score_gpt":0.2810470987044051,"score_spread":0.22165839684428518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394208231","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.071757816,0.0012884315,0.9232323,0.0008310841,0.00030624095,0.00016163287,0.00088277046,0.00068379607,0.0008558357],"genre_scores_gemma":[0.74914384,0.0013528143,0.2386005,0.00039884966,0.0011664679,0.00052113127,0.0042953766,0.00012910337,0.0043919096],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986588,0.00048479481,0.00010791773,0.00039162335,0.00019153766,0.0001653512],"domain_scores_gemma":[0.9964378,0.0024490715,0.00041077525,0.0001577021,0.00044992162,0.00009472812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002262162,0.0012216915,0.00096482056,0.0011876373,0.0005723501,0.0009072125,0.0010872425,0.0012053623,0.0011493448],"category_scores_gemma":[0.0057030027,0.00046964549,0.0017078724,0.0012564785,0.00060036185,0.0011161121,0.00068545644,0.001865599,0.00081965124],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001104846,0.00054629357,0.015779158,0.00048961455,0.0003437647,0.00065653614,0.0014827129,0.5730809,0.01639424,0.014632316,0.014247829,0.36124176],"study_design_scores_gemma":[0.000013730135,0.00003497273,0.0009929278,0.000010539616,0.0000142245135,0.000030114672,0.000037387963,0.9951421,0.00035531403,0.0027294313,0.0006269817,0.000012235756],"about_ca_topic_score_codex":0.0140788825,"about_ca_topic_score_gemma":0.013942351,"teacher_disagreement_score":0.0140788825,"about_ca_system_score_codex":0.0006584109,"about_ca_system_score_gemma":0.00092939293,"threshold_uncertainty_score":0.027993858},"labels":[],"label_agreement":null},{"id":"W4394317164","doi":"10.6084/m9.figshare.21135946.v2","title":"Fuirst et al. RSPB Dataset 3","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.08057569017123814,"score_gpt":0.3097765894338298,"score_spread":0.22920089926259166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394317164","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003329353,0.000084711064,0.00009590832,0.000046510046,0.000019765432,0.000013030705,0.9983328,0.00035692693,0.0007173894],"genre_scores_gemma":[0.0008839832,0.00005313138,0.00042999972,0.000033650416,0.0000047079943,0.000083582854,0.9976921,0.00006788644,0.00075110444],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993642,0.000089386296,0.000056406858,0.0001933744,0.00017411275,0.00012270507],"domain_scores_gemma":[0.9977913,0.0007487845,0.00011169839,0.00040839225,0.0007547109,0.00018513456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007143582,0.0016811073,0.00102238,0.0030205573,0.0012255377,0.0017027273,0.0022430238,0.0016527482,0.045343928],"category_scores_gemma":[0.0063927383,0.0004907135,0.0013692805,0.0045388583,0.00041662992,0.0008034828,0.001288979,0.0013854404,0.037518743],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000065759334,0.000027580236,0.0019400034,0.000599939,0.000062300256,0.000022263574,0.00004795783,0.0008053384,0.00014705567,0.00057205075,0.9921802,0.0035296006],"study_design_scores_gemma":[0.0002448355,0.000021106733,0.01134251,0.0003031341,0.00007966889,0.000061628554,0.00017145135,0.0016719882,0.00038106696,0.0013832086,0.984294,0.000045522636],"about_ca_topic_score_codex":0.2231367,"about_ca_topic_score_gemma":0.362496,"teacher_disagreement_score":0.2231367,"about_ca_system_score_codex":0.001983825,"about_ca_system_score_gemma":0.0045817364,"threshold_uncertainty_score":0.44367582},"labels":[],"label_agreement":null},{"id":"W4394352598","doi":"10.6084/m9.figshare.13368924","title":"Topic Modeling on Triage Notes With Semiorthogonal Nonnegative Matrix Factorization","year":2021,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Factorization; Triage; Matrix (chemical analysis); Matrix decomposition; Algebra over a field; Computer science; Mathematics; Algorithm; Physics; Pure mathematics; Psychology","score_opus":0.06542650446239814,"score_gpt":0.2910039415851521,"score_spread":0.22557743712275394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394352598","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07480161,0.0011363485,0.92016906,0.0008115836,0.00028892793,0.00017246207,0.001012415,0.0005300031,0.0010775098],"genre_scores_gemma":[0.74921244,0.0014141048,0.23780198,0.00036563122,0.001028381,0.00055721536,0.0044876477,0.00014284784,0.004989807],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986355,0.0005213964,0.00010179937,0.000399021,0.0001736021,0.00016879886],"domain_scores_gemma":[0.996594,0.0023376879,0.0004022882,0.0001728541,0.00040009373,0.00009304161],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021874697,0.0012207662,0.00090569153,0.001252511,0.00056242006,0.0010502912,0.0010380266,0.001102314,0.0013673055],"category_scores_gemma":[0.0059974976,0.00047891543,0.0016726427,0.0013411535,0.00061805215,0.0013024508,0.0007373653,0.0017274834,0.000802945],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010855552,0.0005125392,0.014775462,0.0004560546,0.0003115057,0.00065840647,0.0015597355,0.6023784,0.016335301,0.025928084,0.013472221,0.32252672],"study_design_scores_gemma":[0.0000133099775,0.000030480793,0.00083025644,0.000010597138,0.000014731537,0.000029075436,0.000039530634,0.9937662,0.0003350458,0.0042753285,0.0006420772,0.000013300004],"about_ca_topic_score_codex":0.013166824,"about_ca_topic_score_gemma":0.012473513,"teacher_disagreement_score":0.013166824,"about_ca_system_score_codex":0.0007037487,"about_ca_system_score_gemma":0.0008848697,"threshold_uncertainty_score":0.026180327},"labels":[],"label_agreement":null},{"id":"W4394456826","doi":"10.6084/m9.figshare.21135946","title":"Fuirst et al. RSPB Dataset 3","year":2023,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.10801664498180764,"score_gpt":0.3278010660869608,"score_spread":0.21978442110515317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394456826","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003249519,0.00008402869,0.00009448049,0.000046766105,0.000019808078,0.0000128096735,0.9983462,0.00034933526,0.00072161766],"genre_scores_gemma":[0.0008700662,0.000053465676,0.00042506334,0.000033585136,0.0000047370286,0.00008211643,0.9977034,0.00006774402,0.000759819],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99935573,0.00009117546,0.00005758014,0.00019310684,0.0001776117,0.00012477215],"domain_scores_gemma":[0.99772805,0.00076528406,0.000113632435,0.00041644849,0.0007877483,0.00018887327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007225668,0.0016572878,0.0010212673,0.003037994,0.0012258149,0.001706447,0.002236132,0.0016287853,0.04606974],"category_scores_gemma":[0.006561299,0.0004926968,0.0013809648,0.004590936,0.00041747454,0.0008046638,0.0012929469,0.0013849958,0.038124714],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006502977,0.000026326903,0.0018912179,0.00058575266,0.000060973813,0.000021299873,0.00004693654,0.0007849136,0.00014020214,0.0005647082,0.99231905,0.0034936105],"study_design_scores_gemma":[0.0002452608,0.00002082684,0.011273948,0.00030778855,0.00007950926,0.000060124043,0.0001712082,0.001646912,0.00037371783,0.0013992604,0.98437655,0.000044891814],"about_ca_topic_score_codex":0.2307467,"about_ca_topic_score_gemma":0.37108248,"teacher_disagreement_score":0.2307467,"about_ca_system_score_codex":0.0020234461,"about_ca_system_score_gemma":0.0047124433,"threshold_uncertainty_score":0.45880717},"labels":[],"label_agreement":null},{"id":"W4394478980","doi":"10.6084/m9.figshare.21135946.v1","title":"Fuirst et al. RSPB Dataset 3","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.08057569017123814,"score_gpt":0.3097765894338298,"score_spread":0.22920089926259166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394478980","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003329353,0.000084711064,0.00009590832,0.000046510046,0.000019765432,0.000013030705,0.9983328,0.00035692693,0.0007173894],"genre_scores_gemma":[0.0008839832,0.00005313138,0.00042999972,0.000033650416,0.0000047079943,0.000083582854,0.9976921,0.00006788644,0.00075110444],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993642,0.000089386296,0.000056406858,0.0001933744,0.00017411275,0.00012270507],"domain_scores_gemma":[0.9977913,0.0007487845,0.00011169839,0.00040839225,0.0007547109,0.00018513456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007143582,0.0016811073,0.00102238,0.0030205573,0.0012255377,0.0017027273,0.0022430238,0.0016527482,0.045343928],"category_scores_gemma":[0.0063927383,0.0004907135,0.0013692805,0.0045388583,0.00041662992,0.0008034828,0.001288979,0.0013854404,0.037518743],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000065759334,0.000027580236,0.0019400034,0.000599939,0.000062300256,0.000022263574,0.00004795783,0.0008053384,0.00014705567,0.00057205075,0.9921802,0.0035296006],"study_design_scores_gemma":[0.0002448355,0.000021106733,0.01134251,0.0003031341,0.00007966889,0.000061628554,0.00017145135,0.0016719882,0.00038106696,0.0013832086,0.984294,0.000045522636],"about_ca_topic_score_codex":0.2231367,"about_ca_topic_score_gemma":0.362496,"teacher_disagreement_score":0.2231367,"about_ca_system_score_codex":0.001983825,"about_ca_system_score_gemma":0.0045817364,"threshold_uncertainty_score":0.44367582},"labels":[],"label_agreement":null},{"id":"W4394563023","doi":"10.6084/m9.figshare.14040183","title":"Fast Computation of Latent Correlations","year":2021,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computation; Statistical physics; Computer science; Algorithm; Physics","score_opus":0.06741901773838396,"score_gpt":0.2798929596676905,"score_spread":0.21247394192930655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394563023","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035470247,0.0031305125,0.80963415,0.0011304572,0.0002769993,0.00034492763,0.098652005,0.045868594,0.005492191],"genre_scores_gemma":[0.17973663,0.0011716117,0.6000358,0.00044204653,0.00015826603,0.0010685979,0.20932215,0.0030899036,0.0049750637],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978021,0.0007478599,0.00011871382,0.00071577966,0.00037626675,0.00023935975],"domain_scores_gemma":[0.9967007,0.0013968004,0.00017797684,0.0010900918,0.00052300777,0.000111407564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034514696,0.0020032427,0.0019622543,0.0034023926,0.0009728873,0.0026913723,0.0024643615,0.0014731031,0.008901383],"category_scores_gemma":[0.017788837,0.0009676675,0.0028602337,0.0054929047,0.0005317027,0.0032651857,0.0034161245,0.0027906308,0.009359095],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008100933,0.00029675322,0.022517236,0.0013387105,0.0008548595,0.0003033053,0.00062933384,0.2028534,0.0063787447,0.040639415,0.3246529,0.39872524],"study_design_scores_gemma":[0.00018126202,0.00005151022,0.004165275,0.000087300876,0.000103089966,0.0001940865,0.00017209713,0.8740269,0.0035362705,0.067370065,0.050039236,0.000072802366],"about_ca_topic_score_codex":0.013288079,"about_ca_topic_score_gemma":0.024775088,"teacher_disagreement_score":0.013288079,"about_ca_system_score_codex":0.0014228142,"about_ca_system_score_gemma":0.0033564474,"threshold_uncertainty_score":0.029778123},"labels":[],"label_agreement":null},{"id":"W4394580931","doi":"10.35542/osf.io/6d8tj","title":"A Review of Automatic Item Generation Techniques Leveraging Large Language Models","year":2024,"lang":"en","type":"review","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Language model; Data science; Artificial intelligence","score_opus":0.10258415905718281,"score_gpt":0.3687607373528912,"score_spread":0.2661765782957084,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394580931","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030925907,0.90321755,0.08433232,0.0018204655,0.0003782375,0.0007929968,0.0014612293,0.0011265764,0.0037779105],"genre_scores_gemma":[0.02648623,0.73009324,0.23440564,0.0014835671,0.0006391485,0.0019350938,0.0033262172,0.00033961824,0.0012912314],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98615605,0.007586812,0.0022544386,0.0012452757,0.0025954614,0.00016203389],"domain_scores_gemma":[0.89374065,0.09393409,0.003161629,0.002427419,0.0064790617,0.0002570976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021961074,0.002706224,0.003200324,0.0132543985,0.0006696316,0.002864552,0.003990593,0.0014997415,0.0057181264],"category_scores_gemma":[0.07223466,0.001505438,0.004238543,0.012017078,0.00097401364,0.004382811,0.001447397,0.001801185,0.003769623],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010862174,0.000121413206,0.0017231021,0.045397736,0.0007277231,0.000111566216,0.00041243696,0.001944195,0.0006736085,0.0025456853,0.009829079,0.9364049],"study_design_scores_gemma":[0.000443932,0.0015929793,0.045212638,0.1622936,0.009971125,0.0033552165,0.0023165822,0.04289169,0.009801196,0.037084438,0.6841079,0.000928692],"about_ca_topic_score_codex":0.005109525,"about_ca_topic_score_gemma":0.007448732,"teacher_disagreement_score":0.021961074,"about_ca_system_score_codex":0.0014739641,"about_ca_system_score_gemma":0.004540528,"threshold_uncertainty_score":0.11614269},"labels":[],"label_agreement":null},{"id":"W4394593186","doi":"10.1109/wacv57701.2024.00290","title":"Synthesizing Coherent Story with Auto-Regressive Latent Diffusion Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Diffusion; Autoregressive model; Artificial intelligence; Natural language processing; Econometrics; Mathematics; Physics; Thermodynamics","score_opus":0.027657368657227847,"score_gpt":0.2325673091157755,"score_spread":0.20490994045854766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394593186","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03398109,0.00089463306,0.9573874,0.0002905772,0.00008017706,0.000094343726,0.00062319083,0.003944339,0.0027043289],"genre_scores_gemma":[0.48673803,0.0008763335,0.4984732,0.00025818514,0.00010095588,0.0002370283,0.0031126046,0.00085283536,0.009350796],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99967635,0.00010029228,0.000015368973,0.00012293014,0.0000613196,0.000023733524],"domain_scores_gemma":[0.9990404,0.00060514,0.00007959333,0.00013940613,0.00008764874,0.000047753605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007245091,0.00095812755,0.00057794654,0.0006776839,0.00022669228,0.000976984,0.0008769117,0.00083589234,0.0037154725],"category_scores_gemma":[0.003421699,0.00037565466,0.0010759983,0.00049771543,0.0004013249,0.0015372699,0.0008650033,0.001400467,0.0014211424],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048355598,0.00024786164,0.0015506853,0.0005943094,0.0001938632,0.0002760966,0.0004332658,0.38503432,0.047320236,0.019106075,0.010902939,0.5338568],"study_design_scores_gemma":[0.000022798089,0.00004324981,0.00018876567,0.000013334111,0.00001549087,0.00005328393,0.000028096712,0.98736656,0.0055283383,0.004852404,0.001874955,0.000012651746],"about_ca_topic_score_codex":0.0026888447,"about_ca_topic_score_gemma":0.0050046085,"teacher_disagreement_score":0.0037154725,"about_ca_system_score_codex":0.0005050826,"about_ca_system_score_gemma":0.0003964869,"threshold_uncertainty_score":0.012429535},"labels":[],"label_agreement":null},{"id":"W4394647506","doi":"10.1162/tacl_a_00670","title":"Scope Ambiguities in Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; National Research Council Canada; Mila - Quebec Artificial Intelligence Institute","funders":"Fonds de Recherche du Québec-Société et Culture","keywords":"Scope (computer science); Computer science; Linguistics; Cognitive science; Psychology; Programming language; Philosophy","score_opus":0.024324775830183193,"score_gpt":0.2884349640007543,"score_spread":0.2641101881705711,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394647506","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32422498,0.001955675,0.66450864,0.0017971792,0.00008947296,0.00012172361,0.0010379842,0.0018619086,0.0044024996],"genre_scores_gemma":[0.9528204,0.00029830134,0.04485708,0.00017806866,0.00008096759,0.000067724475,0.0010065337,0.0001398449,0.0005511562],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9911237,0.006212833,0.00039819733,0.0011724443,0.00089714874,0.00019580295],"domain_scores_gemma":[0.93122303,0.06177311,0.0029612384,0.002255571,0.0013806575,0.00040644905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012148114,0.00095815846,0.0011454097,0.0024156775,0.0008754414,0.0032000912,0.0011766035,0.0011723896,0.001467629],"category_scores_gemma":[0.05126545,0.00070427475,0.0012645321,0.0022815946,0.001747151,0.004489352,0.0016169407,0.0026255946,0.00036859524],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006892596,0.00019534885,0.027911369,0.00051724777,0.0004694706,0.0010804855,0.0046810526,0.7596234,0.004019022,0.10697125,0.0073307687,0.08651127],"study_design_scores_gemma":[0.000020365314,0.000021587452,0.0014829377,0.000036006015,0.000026209651,0.000076453915,0.00023191486,0.90921956,0.00057952484,0.087237135,0.001040671,0.000027618986],"about_ca_topic_score_codex":0.007378241,"about_ca_topic_score_gemma":0.00832256,"teacher_disagreement_score":0.012148114,"about_ca_system_score_codex":0.0015280144,"about_ca_system_score_gemma":0.0009840496,"threshold_uncertainty_score":0.06424612},"labels":[],"label_agreement":null},{"id":"W4394770144","doi":"10.1162/dint_a_00251","title":"LLaMA-LoRA Neural Prompt Engineering: A Deep Tuning Framework for Automatically Generating Chinese Text Logical Reasoning Thinking Chains","year":2024,"lang":"en","type":"article","venue":"Data Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Language model; Artificial intelligence; Natural language processing; Comprehension; Benchmark (surveying); Inference; Logical reasoning; Question answering; Programming language","score_opus":0.05625293307651095,"score_gpt":0.3232568566701648,"score_spread":0.2670039235936538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394770144","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041934106,0.0004320957,0.9411545,0.0006254954,0.000104235914,0.00019129328,0.00050737365,0.012183863,0.0028671413],"genre_scores_gemma":[0.58964795,0.00027666788,0.4003469,0.0006659911,0.00008256363,0.0005224294,0.0015402756,0.00039283754,0.006524399],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999645,0.00010588715,0.000023040406,0.00012276882,0.000059103422,0.000044286393],"domain_scores_gemma":[0.99908113,0.00042789945,0.00006243294,0.00010005664,0.00025152418,0.00007690626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012673126,0.00085335626,0.0006103454,0.00067616755,0.00043698496,0.0010040447,0.0019565867,0.0011357474,0.0056091375],"category_scores_gemma":[0.0043233414,0.00045501717,0.0008995861,0.00047694234,0.00054305646,0.0017817894,0.0013696626,0.0023642406,0.0014501584],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031793953,0.00040264946,0.0028214839,0.0002496094,0.000097198375,0.00020759259,0.00030698662,0.3964435,0.012275206,0.01430817,0.010943519,0.56162614],"study_design_scores_gemma":[0.0000131040515,0.00003287789,0.00008988483,0.000007178894,0.000008090213,0.000009098275,0.000013125689,0.9935644,0.001065106,0.00444924,0.0007424815,0.0000053201625],"about_ca_topic_score_codex":0.008054056,"about_ca_topic_score_gemma":0.013599532,"teacher_disagreement_score":0.008054056,"about_ca_system_score_codex":0.0011374771,"about_ca_system_score_gemma":0.0017196464,"threshold_uncertainty_score":0.018764377},"labels":[],"label_agreement":null},{"id":"W4394877201","doi":"10.1145/3658673","title":"CGKPN: Cross-Graph Knowledge Propagation Network with Adaptive Connection for Reasoning-Based Machine Reading Comprehension","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Intelligent Systems and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Central China Normal University; Fundamental Research Funds for the Central Universities; Natural Science Foundation of Hubei Province; National Natural Science Foundation of China","keywords":"Computer science; Artificial intelligence; Theoretical computer science; Graph; Comprehension; Semantic memory; Natural language processing; Cognition; Programming language","score_opus":0.025823321524726585,"score_gpt":0.27844403678114726,"score_spread":0.25262071525642066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394877201","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03425945,0.0016406422,0.9495586,0.001099987,0.00023354423,0.00024160341,0.0014135574,0.0074452558,0.0041072164],"genre_scores_gemma":[0.66465914,0.0014844444,0.31702533,0.0011198102,0.00019500087,0.0005999805,0.0051270314,0.0004806328,0.0093086865],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944836,0.00011760519,0.000026226378,0.000249347,0.00010553256,0.000052883435],"domain_scores_gemma":[0.998995,0.0005229631,0.00007980283,0.000139445,0.00020828148,0.00005453458],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009164861,0.0016164444,0.00084370933,0.0014291737,0.0006356845,0.001049706,0.003587299,0.0019791534,0.0033249552],"category_scores_gemma":[0.005005181,0.0007322053,0.0011631348,0.0014824175,0.00082172756,0.002946901,0.001610421,0.0026598384,0.001014935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030755828,0.0002813931,0.0025428245,0.00032015698,0.00023321589,0.00037004016,0.0002623187,0.6132885,0.0049975007,0.011671661,0.017129669,0.34859508],"study_design_scores_gemma":[0.000011613423,0.000024380033,0.00021444868,0.000009938736,0.000025604133,0.00003622326,0.000012704906,0.9901451,0.00061140896,0.007773961,0.001125675,0.000008962654],"about_ca_topic_score_codex":0.023227792,"about_ca_topic_score_gemma":0.024952449,"teacher_disagreement_score":0.023227792,"about_ca_system_score_codex":0.0018780553,"about_ca_system_score_gemma":0.0012774561,"threshold_uncertainty_score":0.046185136},"labels":[],"label_agreement":null},{"id":"W4394877296","doi":"10.3390/biomedinformatics4020062","title":"Recent Advances in Large Language Models for Healthcare","year":2024,"lang":"en","type":"article","venue":"BioMedInformatics","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"New Brunswick Innovation Foundation; Fondation de la recherche en santé du Nouveau-Brunswick","keywords":"Field (mathematics); Health care; Variety (cybernetics); Computer science; Domain (mathematical analysis); Data science; Medical care; Management science; Risk analysis (engineering); Medicine; Political science; Artificial intelligence; Engineering","score_opus":0.03239555851291891,"score_gpt":0.3223995153578027,"score_spread":0.2900039568448838,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394877296","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005617961,0.095470145,0.84847546,0.03251505,0.0015536308,0.00012019045,0.002006729,0.0035830813,0.010657728],"genre_scores_gemma":[0.23406215,0.16129039,0.5610522,0.009658058,0.009809866,0.0006085376,0.008470075,0.0016769076,0.013371744],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965701,0.0019124355,0.0002464075,0.00053912244,0.0006284353,0.00010347578],"domain_scores_gemma":[0.9826533,0.014012199,0.00045694478,0.001278654,0.0013137532,0.00028509248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065315063,0.0012282119,0.00142322,0.0025473384,0.00043648103,0.003739378,0.0022932487,0.0018149553,0.006289649],"category_scores_gemma":[0.023785967,0.0008332474,0.0018412839,0.0033680336,0.0012436267,0.0051518898,0.0022234388,0.0040043667,0.0038154891],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027635336,0.0001535384,0.003186113,0.0018685204,0.0004398314,0.0002617349,0.00046611464,0.13269912,0.0016832313,0.15962514,0.054116923,0.6452234],"study_design_scores_gemma":[0.00003472249,0.00007121005,0.0009787729,0.00050715246,0.00012872653,0.00025540168,0.0001070912,0.59254974,0.0010727084,0.2542858,0.1499074,0.00010124156],"about_ca_topic_score_codex":0.0073100645,"about_ca_topic_score_gemma":0.0059049972,"teacher_disagreement_score":0.0073100645,"about_ca_system_score_codex":0.0025859,"about_ca_system_score_gemma":0.002598318,"threshold_uncertainty_score":0.034542322},"labels":[],"label_agreement":null},{"id":"W4394905126","doi":"","title":"Systèmes de questions-réponses interactifs à grande échelle","year":2022,"lang":"en","type":"dissertation","venue":"theses.fr (ABES)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Scale (ratio); Computer science; Information retrieval; Data science; Geography; Cartography","score_opus":0.018099920187354922,"score_gpt":0.29021277501133447,"score_spread":0.27211285482397957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394905126","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02202101,0.0017232007,0.72846234,0.0019065844,0.00041689147,0.0015996546,0.004981189,0.22104467,0.0178444],"genre_scores_gemma":[0.26922196,0.0018049757,0.6416636,0.0036490138,0.0008046548,0.004510157,0.022816647,0.008846684,0.04668225],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9892524,0.0040618833,0.0009798277,0.0030419286,0.002263951,0.0003999779],"domain_scores_gemma":[0.9831632,0.01150923,0.0006465403,0.0021279603,0.00193417,0.0006189615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009154394,0.003432304,0.0031240613,0.0030604268,0.0021609128,0.008615438,0.0049457583,0.0052704657,0.0295051],"category_scores_gemma":[0.029939018,0.0019002452,0.0025447793,0.0015790405,0.0023536694,0.0112511255,0.005835284,0.003675445,0.018775554],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0065905806,0.0020450251,0.010728555,0.0043552998,0.0012102567,0.0030787962,0.011488133,0.024957119,0.036789004,0.10477623,0.12649302,0.667488],"study_design_scores_gemma":[0.0011514806,0.0008735979,0.006376539,0.000786929,0.00041454958,0.0017619265,0.002574062,0.34126616,0.052793127,0.1251483,0.46616608,0.0006872448],"about_ca_topic_score_codex":0.004768406,"about_ca_topic_score_gemma":0.003239013,"teacher_disagreement_score":0.0295051,"about_ca_system_score_codex":0.0016319783,"about_ca_system_score_gemma":0.0012815136,"threshold_uncertainty_score":0.0987044},"labels":[],"label_agreement":null},{"id":"W4395680939","doi":"10.2196/55798","title":"Using the Natural Language Processing System Medical Named Entity Recognition-Japanese to Analyze Pharmaceutical Care Records: Natural Language Processing Analysis","year":2024,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Natural language processing; Computer science; Natural (archaeology); Named-entity recognition; Natural language; Artificial intelligence; Information retrieval; Engineering; History","score_opus":0.0672794325507549,"score_gpt":0.45641476647827744,"score_spread":0.38913533392752253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395680939","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49901858,0.0034010233,0.4799186,0.0016493694,0.0001601729,0.0013597141,0.006475429,0.0033238058,0.0046931524],"genre_scores_gemma":[0.5374439,0.0014229773,0.45149276,0.00031813697,0.00012115422,0.00050108385,0.006939397,0.00007471487,0.001685931],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9970649,0.0009042566,0.0007001868,0.0006659712,0.00057739945,0.00008720771],"domain_scores_gemma":[0.99464864,0.0028911352,0.0007355992,0.0003713865,0.0012643541,0.00008895231],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043128715,0.0006945149,0.00054616295,0.003724443,0.00053600065,0.0014081502,0.00048613304,0.00053158234,0.0008901826],"category_scores_gemma":[0.0086691845,0.00020090873,0.00087227364,0.002717326,0.00038446794,0.00201617,0.0007376918,0.00043417583,0.00050844916],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006000077,0.00047361242,0.13404474,0.0022102003,0.0005352502,0.0023186372,0.0031028034,0.012495614,0.09125559,0.0027623684,0.0069362903,0.743265],"study_design_scores_gemma":[0.00019541556,0.0010454765,0.31292883,0.00037630103,0.0016127636,0.0049444735,0.0033603106,0.44911116,0.1760049,0.008813586,0.04114443,0.00046231024],"about_ca_topic_score_codex":0.0044962238,"about_ca_topic_score_gemma":0.004179181,"teacher_disagreement_score":0.0044962238,"about_ca_system_score_codex":0.00066985365,"about_ca_system_score_gemma":0.0012135659,"threshold_uncertainty_score":0.02280891},"labels":[],"label_agreement":null},{"id":"W4396624257","doi":"10.48550/arxiv.2405.00492","title":"Is Temperature the Creativity Parameter of Large Language Models?","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Creativity; Linguistics; Psychology; Philosophy; Social psychology","score_opus":0.06541205719938131,"score_gpt":0.21192076922184044,"score_spread":0.14650871202245913,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396624257","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67587495,0.0019016868,0.3043953,0.0057300827,0.00018552282,0.00019628613,0.0005876046,0.00084849604,0.010279995],"genre_scores_gemma":[0.9801788,0.00028010516,0.018034384,0.0003137421,0.0001624994,0.00014969303,0.00028377937,0.00017162849,0.0004254294],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98240995,0.012360101,0.00076212804,0.002632034,0.0012708515,0.00056486734],"domain_scores_gemma":[0.7104347,0.2457337,0.01571378,0.021836005,0.0037794241,0.002502483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02375885,0.00094614987,0.001275478,0.0010446123,0.0009183415,0.004561378,0.0012860548,0.0017626028,0.0026870603],"category_scores_gemma":[0.2211371,0.00092961965,0.0015724591,0.0011082714,0.0033509282,0.0087585095,0.002303243,0.0043638465,0.0006605368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044877757,0.0008484286,0.34855738,0.0013721314,0.0030957458,0.0012033349,0.010512203,0.27426848,0.01799512,0.192571,0.007598553,0.13748993],"study_design_scores_gemma":[0.00043451812,0.00058884645,0.054813482,0.00027954843,0.0004770543,0.0009544263,0.0019126638,0.48711994,0.00999767,0.43897894,0.0040994734,0.00034346044],"about_ca_topic_score_codex":0.0012333873,"about_ca_topic_score_gemma":0.0012295382,"teacher_disagreement_score":0.02375885,"about_ca_system_score_codex":0.0011146202,"about_ca_system_score_gemma":0.0009780779,"threshold_uncertainty_score":0.12565029},"labels":[],"label_agreement":null},{"id":"W4396674157","doi":"10.32920/25761534","title":"Methods and Datasets for Feature-based query refinement collection","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Information retrieval; Query expansion; Web query classification; Web search query; Closeness; Sargable; Query optimization; Query language; Term (time); Vocabulary; RDF query language; Search engine; Mathematics","score_opus":0.05748325247809054,"score_gpt":0.3842360066330495,"score_spread":0.32675275415495897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396674157","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037036054,0.004895042,0.1410768,0.0019875949,0.0016891293,0.01558626,0.7361537,0.034770068,0.02680528],"genre_scores_gemma":[0.03738899,0.0010778925,0.16362287,0.00072715804,0.00014882172,0.012141959,0.77644986,0.0006743008,0.007768186],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9950877,0.0009618646,0.00083459524,0.0011010712,0.0015936088,0.00042110455],"domain_scores_gemma":[0.9936864,0.00144122,0.0003314477,0.0023078143,0.0019751117,0.00025795845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042591123,0.0021172531,0.0013040241,0.005541867,0.0015032444,0.0022256211,0.004683947,0.0027697945,0.029534055],"category_scores_gemma":[0.014983666,0.00076912437,0.002711261,0.0056211106,0.000823099,0.0018527261,0.0021806352,0.00300917,0.029483104],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002061649,0.002731843,0.009557469,0.0033883455,0.00031789456,0.00067535293,0.00029543953,0.021126954,0.008087238,0.007745316,0.5943501,0.34966245],"study_design_scores_gemma":[0.0019887504,0.0013239535,0.023528805,0.0007541753,0.00036592796,0.001837033,0.0010948617,0.13990737,0.028103422,0.01991775,0.7807317,0.0004463945],"about_ca_topic_score_codex":0.019391768,"about_ca_topic_score_gemma":0.020944735,"teacher_disagreement_score":0.029534055,"about_ca_system_score_codex":0.0019855176,"about_ca_system_score_gemma":0.0031867097,"threshold_uncertainty_score":0.098801196},"labels":[],"label_agreement":null},{"id":"W4396674458","doi":"10.32920/25761534.v1","title":"Methods and Datasets for Feature-based query refinement collection","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Information retrieval; Query expansion; Web query classification; Web search query; Closeness; Query optimization; Sargable; Query language; Term (time); Vocabulary; RDF query language; Search engine; Mathematics","score_opus":0.05748325247809054,"score_gpt":0.3842360066330495,"score_spread":0.32675275415495897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396674458","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037036054,0.004895042,0.1410768,0.0019875949,0.0016891293,0.01558626,0.7361537,0.034770068,0.02680528],"genre_scores_gemma":[0.03738899,0.0010778925,0.16362287,0.00072715804,0.00014882172,0.012141959,0.77644986,0.0006743008,0.007768186],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9950877,0.0009618646,0.00083459524,0.0011010712,0.0015936088,0.00042110455],"domain_scores_gemma":[0.9936864,0.00144122,0.0003314477,0.0023078143,0.0019751117,0.00025795845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042591123,0.0021172531,0.0013040241,0.005541867,0.0015032444,0.0022256211,0.004683947,0.0027697945,0.029534055],"category_scores_gemma":[0.014983666,0.00076912437,0.002711261,0.0056211106,0.000823099,0.0018527261,0.0021806352,0.00300917,0.029483104],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002061649,0.002731843,0.009557469,0.0033883455,0.00031789456,0.00067535293,0.00029543953,0.021126954,0.008087238,0.007745316,0.5943501,0.34966245],"study_design_scores_gemma":[0.0019887504,0.0013239535,0.023528805,0.0007541753,0.00036592796,0.001837033,0.0010948617,0.13990737,0.028103422,0.01991775,0.7807317,0.0004463945],"about_ca_topic_score_codex":0.019391768,"about_ca_topic_score_gemma":0.020944735,"teacher_disagreement_score":0.029534055,"about_ca_system_score_codex":0.0019855176,"about_ca_system_score_gemma":0.0031867097,"threshold_uncertainty_score":0.098801196},"labels":[],"label_agreement":null},{"id":"W4396722757","doi":"10.1145/3589334.3645481","title":"Metacognitive Retrieval-Augmented Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Universitas Brawijaya","keywords":"Computer science; Metacognition; Natural language processing; Artificial intelligence; Language model; Information retrieval; Cognition; Psychology","score_opus":0.025825323228015694,"score_gpt":0.2830062823510828,"score_spread":0.25718095912306715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396722757","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07285092,0.00059243327,0.9178986,0.00059283106,0.000046665366,0.00013955189,0.00041606667,0.0030427799,0.004420201],"genre_scores_gemma":[0.8031701,0.00040178664,0.19045947,0.00014208336,0.000069847454,0.0003566498,0.00067544484,0.00020046026,0.004524193],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999303,0.0003443338,0.00003613644,0.00014615928,0.000108530614,0.00006187522],"domain_scores_gemma":[0.99538875,0.0034042483,0.0003066609,0.00046018485,0.00031783248,0.00012225355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018385105,0.0007951304,0.0006576947,0.0006618854,0.00031174318,0.0019234465,0.0018197204,0.00089939503,0.0026129198],"category_scores_gemma":[0.007475322,0.00038999048,0.0010488684,0.00047968066,0.00062598387,0.0023316825,0.0010808526,0.0016013331,0.0006495967],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025382475,0.00017394684,0.0018019798,0.00012294232,0.00010766441,0.00012095319,0.00033656863,0.87616426,0.0019602857,0.038617805,0.0016856958,0.07865407],"study_design_scores_gemma":[0.000010364383,0.000012282693,0.00006601056,0.0000030175481,0.000008607145,0.000007477276,0.000006037442,0.99120444,0.00025678024,0.008115798,0.00030517674,0.0000039165852],"about_ca_topic_score_codex":0.005123515,"about_ca_topic_score_gemma":0.008319078,"teacher_disagreement_score":0.005123515,"about_ca_system_score_codex":0.0008909081,"about_ca_system_score_gemma":0.0012043053,"threshold_uncertainty_score":0.0101873875},"labels":[],"label_agreement":null},{"id":"W4396762941","doi":"10.3389/fdata.2024.1295009","title":"Multi-modal recommender system for predicting project manager performance within a competency-based framework","year":2024,"lang":"en","type":"article","venue":"Frontiers in Big Data","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université TÉLUQ","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Recommender system; Task (project management); Process (computing); Collaborative filtering; Context (archaeology); Modalities; Project manager; Information retrieval; Machine learning; Artificial intelligence; Knowledge management; Project management; Engineering","score_opus":0.09354698843037028,"score_gpt":0.2997290621673356,"score_spread":0.2061820737369653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396762941","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45595038,0.00341331,0.52546203,0.0010605602,0.00030908157,0.0004287618,0.0041927435,0.0051139873,0.0040691933],"genre_scores_gemma":[0.8246866,0.00065459864,0.16689382,0.00025300006,0.0001447033,0.00027099313,0.003933956,0.00005147273,0.003110874],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986644,0.00039518694,0.00014398493,0.00043989086,0.00024269798,0.00011392775],"domain_scores_gemma":[0.99779886,0.0011023536,0.00018558248,0.0001694261,0.0006385859,0.00010509847],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023134947,0.0010793081,0.0012705142,0.0023769813,0.00048233682,0.0008168478,0.0011851595,0.0012467459,0.0011725387],"category_scores_gemma":[0.004439414,0.00039042294,0.0011424217,0.0012911175,0.0001595359,0.00083882187,0.00047669056,0.0010242844,0.00104721],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012216442,0.0014867119,0.14728794,0.00064511586,0.001194344,0.0008176724,0.0007100938,0.17120609,0.023758352,0.0016441576,0.016992763,0.63303506],"study_design_scores_gemma":[0.000028562474,0.00022475705,0.018496001,0.00003672338,0.00012357532,0.00016782024,0.00010570339,0.9762645,0.0024102153,0.00042400474,0.0016691881,0.00004893559],"about_ca_topic_score_codex":0.028082985,"about_ca_topic_score_gemma":0.033080693,"teacher_disagreement_score":0.028082985,"about_ca_system_score_codex":0.00059223507,"about_ca_system_score_gemma":0.0007032094,"threshold_uncertainty_score":0.05583906},"labels":[],"label_agreement":null},{"id":"W4396775553","doi":"10.2196/54633","title":"A Reliable and Accessible Caregiving Language Model (CaLM) to Support Tools for Caregivers: Development and Evaluation Study","year":2024,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Institute on Disability, Independent Living, and Rehabilitation Research; Administration for Community Living","keywords":"Computer science; Quality (philosophy); Set (abstract data type); Reliability (semiconductor); Family caregivers; Psychology; Medicine; Nursing","score_opus":0.1869762067679215,"score_gpt":0.47128561755562876,"score_spread":0.2843094107877072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396775553","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89038306,0.0015645772,0.0910154,0.00058999605,0.00019141201,0.0029730352,0.0017754639,0.0063849934,0.005122097],"genre_scores_gemma":[0.7811347,0.001113001,0.2048987,0.0003062763,0.00003200705,0.002337549,0.0065691397,0.00027997233,0.0033286286],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998539,0.00070205843,0.00012972923,0.00019699754,0.0003451302,0.00008712608],"domain_scores_gemma":[0.99497885,0.003093197,0.0002089629,0.00038788098,0.0010228984,0.00030819137],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043544946,0.0013038247,0.00053837185,0.00065522565,0.00028346633,0.00078418717,0.00224849,0.0009572584,0.0025234448],"category_scores_gemma":[0.0144320745,0.00037729766,0.00062219816,0.00031441607,0.00044456465,0.0018113436,0.0014596699,0.0014706888,0.0007947516],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021403874,0.0062024216,0.02164224,0.0036728806,0.00044673594,0.0014119667,0.003186681,0.12642582,0.018828599,0.0034221122,0.02275658,0.78986365],"study_design_scores_gemma":[0.0011195164,0.0059206267,0.010638373,0.00044951463,0.00042923226,0.0006720058,0.0014508981,0.936548,0.021352544,0.0014906329,0.019787535,0.00014115674],"about_ca_topic_score_codex":0.0076429914,"about_ca_topic_score_gemma":0.0054351673,"teacher_disagreement_score":0.0076429914,"about_ca_system_score_codex":0.0011684146,"about_ca_system_score_gemma":0.001987497,"threshold_uncertainty_score":0.02302903},"labels":[],"label_agreement":null},{"id":"W4396802056","doi":"10.1145/3630106.3658941","title":"\"I'm Not Sure, But...\": Examining the Impact of Large Language Models' Uncertainty Expression on User Reliance and Trust","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":105,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"Microsoft Research; Universitas Brawijaya; National Science Foundation","keywords":"Scale (ratio); Task (project management); Natural experiment; Expression (computer science); Psychology; Perspective (graphical); Computer science; Social psychology; Cognitive psychology; Artificial intelligence; Medicine; Engineering","score_opus":0.04877756425759006,"score_gpt":0.3127205170671762,"score_spread":0.26394295280958613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396802056","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.995145,0.00007646298,0.003799029,0.00015740065,0.0000054180855,0.000029258068,0.000027472674,0.00006194778,0.0006978966],"genre_scores_gemma":[0.9963051,0.000044818287,0.0032784431,0.00011224417,0.000006689273,0.000048016238,0.00003353212,0.000020693797,0.00015045778],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9894874,0.008306937,0.00047609463,0.00064276153,0.0008533554,0.00023337918],"domain_scores_gemma":[0.8086087,0.16712795,0.013923618,0.005849748,0.003345453,0.0011445514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013858043,0.0005542659,0.0003884805,0.00038012056,0.0005018368,0.0020302846,0.0005867003,0.00080683175,0.0013864494],"category_scores_gemma":[0.10733189,0.00040328773,0.0004281413,0.00024279239,0.0012135404,0.0031415827,0.0012591921,0.001123167,0.0003069648],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011876996,0.0031351664,0.41817474,0.0035501502,0.00082034105,0.0012897311,0.18076508,0.010971207,0.1334572,0.006609154,0.002577114,0.22677314],"study_design_scores_gemma":[0.0010331221,0.016031943,0.6440717,0.0011929674,0.0018954505,0.002512923,0.06073117,0.14417364,0.089709446,0.021380682,0.016138487,0.0011285259],"about_ca_topic_score_codex":0.0014653542,"about_ca_topic_score_gemma":0.001450392,"teacher_disagreement_score":0.013858043,"about_ca_system_score_codex":0.00058457244,"about_ca_system_score_gemma":0.00055405544,"threshold_uncertainty_score":0.073289156},"labels":[],"label_agreement":null},{"id":"W4396819963","doi":"10.1145/3626772.3657864","title":"Ranked List Truncation for Large Language Model-based Re-Ranking","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"HORIZON EUROPE Framework Programme; China Scholarship Council; European Commission; Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Baidu","keywords":"Computer science; Ranking (information retrieval); Truncation (statistics); Information retrieval; Natural language processing; Machine learning","score_opus":0.025416990242419238,"score_gpt":0.29312282429812736,"score_spread":0.26770583405570814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396819963","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02808238,0.0038712395,0.95260334,0.0008830483,0.0003247921,0.00041407437,0.0009254761,0.0085378345,0.0043577245],"genre_scores_gemma":[0.34456825,0.0013455602,0.6342457,0.0008997721,0.0006535103,0.000559243,0.0035560552,0.001834183,0.012337657],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99425215,0.002620071,0.00041250687,0.0008320301,0.0014507556,0.00043253295],"domain_scores_gemma":[0.9818747,0.009655662,0.0011770454,0.0045560035,0.0022314969,0.00050511974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062330933,0.0020464328,0.0029373066,0.0028723942,0.0013546344,0.002873785,0.0032362523,0.0019118957,0.0062911743],"category_scores_gemma":[0.026917651,0.000735982,0.0017715725,0.0031313626,0.0014291186,0.005557876,0.0025875277,0.0027363189,0.0056633954],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076911715,0.00064262416,0.003871964,0.0012364547,0.00035495593,0.00035805782,0.00067416445,0.2279345,0.015529595,0.03474597,0.04158789,0.6722947],"study_design_scores_gemma":[0.000099973986,0.0003327204,0.00068177824,0.000056043376,0.00010196174,0.00025377635,0.00013684084,0.94667476,0.007788318,0.034124613,0.009669571,0.0000795566],"about_ca_topic_score_codex":0.0060910056,"about_ca_topic_score_gemma":0.012788315,"teacher_disagreement_score":0.0062911743,"about_ca_system_score_codex":0.001715147,"about_ca_system_score_gemma":0.0028939066,"threshold_uncertainty_score":0.03296417},"labels":[],"label_agreement":null},{"id":"W4396820721","doi":"10.48550/arxiv.2404.18424","title":"PromptReps: Prompting Large Language Models to Generate Dense and Sparse Representations for Zero-Shot Document Retrieval","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; University of Waterloo","keywords":"Zero (linguistics); Computer science; Shot (pellet); Natural language processing; Artificial intelligence; Information retrieval; Linguistics; Philosophy; Chemistry","score_opus":0.1266048190329702,"score_gpt":0.25148648447012134,"score_spread":0.12488166543715115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396820721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02412989,0.00063916924,0.9509905,0.00030621854,0.00017353978,0.0001921863,0.00069864,0.021133788,0.0017360906],"genre_scores_gemma":[0.31814885,0.0005598697,0.6587395,0.0007727569,0.0002913824,0.00059724075,0.0061679156,0.0014932707,0.013229285],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990085,0.00034536346,0.000060014827,0.00024469857,0.00024407326,0.000097425844],"domain_scores_gemma":[0.9983187,0.00065512425,0.00011413034,0.0005390784,0.00027887354,0.00009416477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013812752,0.0014599931,0.0012784512,0.0009835463,0.00045778218,0.0011761289,0.0023065798,0.0013590041,0.0051719886],"category_scores_gemma":[0.0055482066,0.00054394745,0.0010516364,0.0008787853,0.00079633447,0.003788663,0.002461748,0.0019699202,0.0045962366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067506987,0.00049915,0.001143482,0.0005517668,0.00017647707,0.00029111386,0.0003646862,0.07871855,0.03943756,0.013163751,0.045049552,0.8199289],"study_design_scores_gemma":[0.00010058083,0.00024145444,0.00029595607,0.000017406137,0.00003227474,0.00014982128,0.00008372565,0.9628185,0.017089091,0.012687176,0.006431402,0.00005262583],"about_ca_topic_score_codex":0.004392797,"about_ca_topic_score_gemma":0.008547316,"teacher_disagreement_score":0.0051719886,"about_ca_system_score_codex":0.0008069597,"about_ca_system_score_gemma":0.0013263728,"threshold_uncertainty_score":0.017302036},"labels":[],"label_agreement":null},{"id":"W4396827151","doi":"10.1145/3613904.3641895","title":"The HaLLMark Effect: Supporting Provenance and Transparent Use of Large Language Models in Writing with Interactive Visualization","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Agency (philosophy); Computer science; Visualization; Control (management); Artificial intelligence; Sociology; Social science","score_opus":0.018483806065146144,"score_gpt":0.30863217937495546,"score_spread":0.2901483733098093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396827151","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08169367,0.0005782403,0.8344019,0.0022139177,0.00029814435,0.00029071284,0.0018227174,0.068962075,0.009738576],"genre_scores_gemma":[0.46921742,0.0003912196,0.51794267,0.00036734127,0.000113716625,0.0003493801,0.0013160437,0.00452764,0.0057745306],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977024,0.0014134243,0.00014742384,0.0002904641,0.00033802018,0.00010817311],"domain_scores_gemma":[0.9785096,0.014593575,0.0010612917,0.0040011914,0.0010002074,0.0008341518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059210495,0.0011249133,0.0005607522,0.0018087752,0.0010386133,0.0046482193,0.0016619286,0.0015866727,0.006669917],"category_scores_gemma":[0.0317663,0.00061734184,0.0010367734,0.0009173047,0.002071636,0.007528046,0.0058935797,0.0022251902,0.0011659558],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035937147,0.0008451228,0.030250981,0.0025479314,0.00034990165,0.0030488262,0.08003781,0.045248076,0.087284364,0.15779668,0.09102321,0.4979734],"study_design_scores_gemma":[0.00075711886,0.00071146956,0.008231998,0.00072873867,0.00030718048,0.0012277327,0.0057173884,0.44037747,0.08582647,0.19767334,0.2578811,0.0005599443],"about_ca_topic_score_codex":0.0028062116,"about_ca_topic_score_gemma":0.004455874,"teacher_disagreement_score":0.006669917,"about_ca_system_score_codex":0.00087235717,"about_ca_system_score_gemma":0.0015542145,"threshold_uncertainty_score":0.031313896},"labels":[],"label_agreement":null},{"id":"W4396828482","doi":"10.1145/3613904.3642462","title":"DirectGPT: A Direct Manipulation Interface to Interact with Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada Research Chairs; University of Toronto; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Undo; Computer science; Syntax; Interface (matter); Programming language; Reuse; Representation (politics); Code (set theory); Software; Human–computer interaction; Artificial intelligence; Operating system; Set (abstract data type)","score_opus":0.021706153306981377,"score_gpt":0.28539954184626976,"score_spread":0.26369338853928836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396828482","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016072676,0.00013760028,0.8440153,0.0002891357,0.00009698966,0.0002628183,0.0018029411,0.12965776,0.0076648123],"genre_scores_gemma":[0.27254087,0.00033030062,0.6685518,0.000678308,0.00008760599,0.0014071371,0.005804064,0.028259782,0.02234019],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991303,0.00032495477,0.00006855646,0.00018160536,0.0002360037,0.000058590816],"domain_scores_gemma":[0.99538785,0.0029921553,0.00015704123,0.0010299139,0.000258832,0.00017428861],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016215498,0.0013183358,0.00052771,0.0006811441,0.00036973853,0.001811091,0.0015853189,0.0011430292,0.029880526],"category_scores_gemma":[0.009630896,0.0006500559,0.0010254566,0.0003176175,0.0007078818,0.003645554,0.004238026,0.0014420355,0.008053534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018792173,0.00072300946,0.007748431,0.0026216593,0.00025826794,0.002408862,0.014938593,0.016411675,0.15328087,0.063907646,0.17288072,0.56294113],"study_design_scores_gemma":[0.00042640208,0.0009361216,0.0050615235,0.0006116605,0.00017752923,0.0024739278,0.0013722931,0.2969672,0.099971384,0.06622766,0.5254035,0.00037090678],"about_ca_topic_score_codex":0.0011530898,"about_ca_topic_score_gemma":0.0019959765,"teacher_disagreement_score":0.029880526,"about_ca_system_score_codex":0.0004310088,"about_ca_system_score_gemma":0.00063323736,"threshold_uncertainty_score":0.09996033},"labels":[],"label_agreement":null},{"id":"W4396831993","doi":"10.1145/3613904.3641960","title":"Human-LLM Collaborative Annotation Through Effective Verification of LLM Labels","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":103,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Annotation; Leverage (statistics); Computer science; Domain (mathematical analysis); Natural language processing; Reliability (semiconductor); Artificial intelligence; Information retrieval; Data science","score_opus":0.019726407796484074,"score_gpt":0.3134079163418643,"score_spread":0.2936815085453802,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396831993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02455704,0.00064814877,0.95597774,0.0027143632,0.00024655348,0.00063598476,0.0009340029,0.009563405,0.004722826],"genre_scores_gemma":[0.2699222,0.00023436498,0.72033983,0.00090748334,0.00020975988,0.00078800187,0.002083711,0.0016297984,0.0038848463],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9082809,0.0605997,0.0034886277,0.014586942,0.01185546,0.0011884308],"domain_scores_gemma":[0.6962627,0.18914741,0.016105698,0.071472466,0.024594808,0.0024169746],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06115983,0.003270684,0.002507101,0.0063278307,0.0040120077,0.006390476,0.0064466298,0.0053567705,0.006224176],"category_scores_gemma":[0.18953283,0.0013764147,0.0026361095,0.0032332188,0.004298394,0.009495041,0.01585355,0.0053836205,0.0043039527],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018207643,0.0009397337,0.019361217,0.002557142,0.00087794353,0.0015450178,0.021189846,0.041478865,0.051235624,0.049719043,0.042913288,0.7663614],"study_design_scores_gemma":[0.00035004976,0.00040861312,0.007202639,0.00086775905,0.00050611375,0.00077220134,0.0047684144,0.6312634,0.07745922,0.17741182,0.09836171,0.0006279985],"about_ca_topic_score_codex":0.006371085,"about_ca_topic_score_gemma":0.011398286,"teacher_disagreement_score":0.93884015,"about_ca_system_score_codex":0.0037087924,"about_ca_system_score_gemma":0.0075569805,"threshold_uncertainty_score":0.323448},"labels":[],"label_agreement":null},{"id":"W4396832250","doi":"10.1145/3613904.3642459","title":"Generative Echo Chamber? Effect of LLM-Powered Search Systems on Diverse Information Seeking","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":98,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Limiting; Echo (communications protocol); Computer science; Polarization (electrochemistry); Keyword search; Psychology; Internet privacy; Computer security; Information retrieval; Engineering","score_opus":0.02109711608346044,"score_gpt":0.27112608044287384,"score_spread":0.2500289643594134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396832250","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98426867,0.0002805105,0.0020982944,0.0007340283,0.000046850786,0.00009252362,0.00012255565,0.00012317115,0.012233414],"genre_scores_gemma":[0.99696356,0.00007484498,0.0013735407,0.00032779452,0.000029627057,0.00010103953,0.000092582304,0.000047530197,0.0009894927],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99440694,0.004207918,0.00023055034,0.00050916366,0.00038602736,0.0002594705],"domain_scores_gemma":[0.8914721,0.09457919,0.004937519,0.00516683,0.0013682813,0.0024761704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071047926,0.00040264238,0.00050660694,0.0004016219,0.0009414644,0.0033369164,0.0005190151,0.0013732467,0.010039755],"category_scores_gemma":[0.07396958,0.0004020034,0.0004252414,0.00028826436,0.001043069,0.0032737432,0.0028058207,0.001684824,0.0012412848],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.040288188,0.009098921,0.30479148,0.004059756,0.0009941657,0.00138788,0.108471446,0.013996692,0.16480479,0.05971395,0.011941847,0.2804508],"study_design_scores_gemma":[0.005366723,0.022132752,0.5597951,0.0015142007,0.0028143993,0.0018093381,0.059514023,0.12744687,0.05938007,0.09443276,0.06496757,0.0008262266],"about_ca_topic_score_codex":0.0011800403,"about_ca_topic_score_gemma":0.00096679805,"teacher_disagreement_score":0.010039755,"about_ca_system_score_codex":0.00066039397,"about_ca_system_score_gemma":0.0006036032,"threshold_uncertainty_score":0.03757423},"labels":[],"label_agreement":null},{"id":"W4396832282","doi":"10.1145/3613904.3641899","title":"ABScribe: Rapid Exploration &amp; Organization of Multiple Writing Variations in Human-AI Co-Writing Tasks using Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":48,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Workflow; Workload; Computer science; Task (project management); Writing process; Process (computing); Variation (astronomy); Human–computer interaction; Perception; Natural language processing; Programming language; Psychology; Database; Engineering; Mathematics education","score_opus":0.06407113873354855,"score_gpt":0.3229418175725084,"score_spread":0.25887067883895987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396832282","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068816505,0.00058057456,0.76785576,0.00059967756,0.0002514738,0.0007715004,0.0035327065,0.14522709,0.012364601],"genre_scores_gemma":[0.22049987,0.00040537625,0.74137884,0.0006590033,0.00009160377,0.0015061368,0.005515189,0.008208209,0.02173569],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9989588,0.00038703627,0.00007814773,0.00028416136,0.00023808687,0.000053798783],"domain_scores_gemma":[0.99287194,0.004838211,0.00024867893,0.0012594912,0.0004645433,0.00031713076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023923898,0.0014444697,0.00063837715,0.00080616545,0.00041788362,0.002104331,0.0021870877,0.0013865015,0.023089169],"category_scores_gemma":[0.013365368,0.0006687858,0.0011834562,0.00038427944,0.00042721073,0.0034122812,0.003675164,0.0011285164,0.0073634167],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022677393,0.0009962952,0.0050918553,0.001992889,0.00015395593,0.00077233487,0.008296999,0.005207684,0.070504315,0.008031059,0.10664531,0.7900396],"study_design_scores_gemma":[0.0016969643,0.0037355635,0.020207921,0.00091515336,0.00026777518,0.0032645676,0.0036180315,0.34191072,0.085197285,0.034293167,0.5040533,0.00083964574],"about_ca_topic_score_codex":0.0007909792,"about_ca_topic_score_gemma":0.002174157,"teacher_disagreement_score":0.023089169,"about_ca_system_score_codex":0.00032735453,"about_ca_system_score_gemma":0.0005257444,"threshold_uncertainty_score":0.077240944},"labels":[],"label_agreement":null},{"id":"W4396843976","doi":"10.1145/3589335.3641299","title":"Information Retrieval Meets Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Information retrieval; Natural language processing; Artificial intelligence; Data science","score_opus":0.014748737637220805,"score_gpt":0.2531784605361799,"score_spread":0.23842972289895906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396843976","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00499544,0.008621421,0.9108686,0.035277564,0.0009948433,0.00025786323,0.0026937264,0.0036672154,0.032623235],"genre_scores_gemma":[0.31619042,0.017021239,0.58074296,0.010855421,0.010963281,0.002446928,0.012549652,0.0034332834,0.04579683],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97503304,0.014906306,0.0014126855,0.0033534807,0.0044530258,0.0008414456],"domain_scores_gemma":[0.8749546,0.101308115,0.0024026325,0.0143918935,0.005628738,0.0013140434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017403701,0.0017936315,0.003808814,0.003295397,0.0025679534,0.011421897,0.0033502036,0.006551488,0.024711832],"category_scores_gemma":[0.09952142,0.0026356208,0.0029510637,0.0037501066,0.0027975475,0.027713975,0.008944652,0.008599915,0.016989592],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000244228,0.0001647794,0.0010469921,0.0010046395,0.00029603942,0.0006855751,0.0007083195,0.023754459,0.0016434883,0.7633619,0.0917307,0.11535891],"study_design_scores_gemma":[0.000071689035,0.000029293395,0.00020017583,0.00008331593,0.00004735085,0.00036123858,0.00013245546,0.16536064,0.00084837776,0.7936571,0.039161187,0.000047149293],"about_ca_topic_score_codex":0.0052911127,"about_ca_topic_score_gemma":0.0056354413,"teacher_disagreement_score":0.024711832,"about_ca_system_score_codex":0.0039442414,"about_ca_system_score_gemma":0.0042686984,"threshold_uncertainty_score":0.09204072},"labels":[],"label_agreement":null},{"id":"W4396882454","doi":"10.48550/arxiv.2405.06563","title":"What Can Natural Language Processing Do for Peer Review?","year":2024,"lang":"en","type":"preprint","venue":"TUbilio (Technical University of Darmstadt)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Office of Naval Research; Bundesministerium für Bildung und Forschung; Deutsche Forschungsgemeinschaft; Natural Sciences and Engineering Research Council of Canada; European Commission; Australian Government; Alberta Machine Intelligence Institute; Canadian Institute for Advanced Research; National Science Foundation","keywords":"Operationalization; Computer science; Pace; Process (computing); Peer review; Artificial intelligence; Technical peer review; Field (mathematics); Aside; Data science; Political science; Linguistics","score_opus":0.023000052290619608,"score_gpt":0.2823599422652428,"score_spread":0.25935988997462317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396882454","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008590429,0.014301423,0.037528414,0.8856276,0.0362703,0.0011338809,0.0006927924,0.002447308,0.021139208],"genre_scores_gemma":[0.06878344,0.06041276,0.36637715,0.33866438,0.104464576,0.009654086,0.0056395647,0.0052723493,0.040731624],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.49665758,0.32808822,0.04271554,0.023977043,0.099848144,0.008713415],"domain_scores_gemma":[0.15434164,0.41506746,0.036957934,0.094570585,0.26961362,0.029448852],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.42656714,0.0025246462,0.0068763527,0.0189651,0.019101407,0.05686799,0.013193916,0.02320452,0.02692041],"category_scores_gemma":[0.7068497,0.0024434286,0.004314145,0.01608667,0.025799248,0.09826195,0.01771663,0.020338185,0.04423434],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013049161,0.00011154473,0.0011330752,0.004149311,0.00016666527,0.00034813996,0.004063141,0.00042223133,0.00048176778,0.036926664,0.7023188,0.2497481],"study_design_scores_gemma":[0.00011569904,0.000058665188,0.0010578478,0.003787339,0.00006314616,0.00017521542,0.0033046412,0.00084349397,0.00035688968,0.10332289,0.8866853,0.00022884588],"about_ca_topic_score_codex":0.0077940505,"about_ca_topic_score_gemma":0.009958006,"teacher_disagreement_score":0.57343286,"about_ca_system_score_codex":0.013956334,"about_ca_system_score_gemma":0.09389085,"threshold_uncertainty_score":0.70714486},"labels":[],"label_agreement":null},{"id":"W4396994480","doi":"10.1681/asn.20213210s1125b","title":"The Quality of Discharge Summaries After AKI","year":2021,"lang":"en","type":"article","venue":"Journal of the American Society of Nephrology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Medicine; Quality (philosophy); Intensive care medicine; Physics","score_opus":0.02536865294790538,"score_gpt":0.29882010960994254,"score_spread":0.2734514566620372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396994480","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96318394,0.011085353,0.0027111308,0.004649686,0.00021221521,0.00019974215,0.015098603,0.00016395061,0.0026954366],"genre_scores_gemma":[0.9907506,0.0018447315,0.00162313,0.00031198587,0.00016382326,0.00006128824,0.00504816,0.000016405627,0.00017992224],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9869108,0.0043869973,0.0044252416,0.0007160163,0.0030798547,0.00048104467],"domain_scores_gemma":[0.81324065,0.041716214,0.120708406,0.0032342472,0.01689275,0.0042077466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007361932,0.0003215714,0.0005674103,0.0037227874,0.00036978855,0.0015015748,0.0006806964,0.0005042797,0.0012267102],"category_scores_gemma":[0.0955693,0.00020534064,0.0007433546,0.0050582723,0.0003537059,0.0018156526,0.0014358945,0.00070724013,0.00019079389],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001813753,0.000027821301,0.98424244,0.00048934773,0.000200625,0.0000996665,0.0004060294,0.00020438737,0.00010680843,0.00009216827,0.001831857,0.01211736],"study_design_scores_gemma":[0.000012799412,0.00014179194,0.99536175,0.00039466476,0.00008966515,0.00030718005,0.0006259177,0.00044428336,0.0002228691,0.00016450507,0.002214462,0.000020066944],"about_ca_topic_score_codex":0.0032555135,"about_ca_topic_score_gemma":0.0029103346,"teacher_disagreement_score":0.007361932,"about_ca_system_score_codex":0.00113613,"about_ca_system_score_gemma":0.0011296425,"threshold_uncertainty_score":0.03893411},"labels":[],"label_agreement":null},{"id":"W4398289211","doi":"10.7910/dvn/eaxeet/o7vrne","title":"GoethePoetryReducedMinusLarge.zip","year":2019,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science","score_opus":0.0322293650370897,"score_gpt":0.2631626910205373,"score_spread":0.2309333259834476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398289211","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002572165,0.00009404646,0.00009122591,0.000099219906,0.000042614654,0.000014939476,0.9971871,0.0011737241,0.0010399551],"genre_scores_gemma":[0.00066498114,0.000069541275,0.00032147494,0.00005296393,0.000014388215,0.00006018422,0.9978149,0.0001500279,0.0008514998],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990632,0.00015202383,0.00009512601,0.00029798108,0.00020367069,0.0001879413],"domain_scores_gemma":[0.9980621,0.0004510958,0.00015859299,0.00064593134,0.00041964228,0.0002626846],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010137729,0.0033156106,0.0013719916,0.0047349464,0.0010462212,0.0030940059,0.0033844782,0.002301116,0.086717],"category_scores_gemma":[0.0057562008,0.00077326474,0.0013135136,0.006455246,0.0006012583,0.0018351715,0.0026654343,0.0015729862,0.14122151],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035737135,0.000017458315,0.0003958886,0.00025355964,0.000012835209,0.00000962833,0.000013410478,0.00012147132,0.00006137613,0.00026075167,0.99706477,0.0017530196],"study_design_scores_gemma":[0.0002874728,0.000028641853,0.002479056,0.00021707655,0.000028700995,0.00006782569,0.00009997389,0.0010140362,0.0006840043,0.0017017148,0.99336207,0.000029491133],"about_ca_topic_score_codex":0.016927835,"about_ca_topic_score_gemma":0.032879937,"teacher_disagreement_score":0.913283,"about_ca_system_score_codex":0.001505527,"about_ca_system_score_gemma":0.0021068528,"threshold_uncertainty_score":0.2900973},"labels":[],"label_agreement":null},{"id":"W4398524366","doi":"10.7910/dvn/ktwyz6/euczd9","title":"simtax_old-checkpoint.py","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal","funders":"","keywords":"Computer science","score_opus":0.027839974531473886,"score_gpt":0.24334158909999085,"score_spread":0.21550161456851696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398524366","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002041808,0.000073036645,0.00013990261,0.00009032206,0.00005108806,0.000014075287,0.9959544,0.0025434364,0.00092961686],"genre_scores_gemma":[0.00068645034,0.00006029198,0.00040033308,0.00007308122,0.000016027074,0.00007560368,0.99749124,0.00036452484,0.00083242136],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999094,0.0001655009,0.000080095604,0.00030892613,0.00017364901,0.0001779349],"domain_scores_gemma":[0.99777126,0.00053203007,0.00014685794,0.00086347543,0.00038417478,0.00030212052],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001448676,0.0033003998,0.0015092661,0.0038307037,0.0009577341,0.0032355,0.0038780498,0.002334173,0.119577914],"category_scores_gemma":[0.006784124,0.000963511,0.00174843,0.0050414633,0.00069354865,0.0019753906,0.0027545323,0.0017982911,0.19378085],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041991774,0.000015826588,0.00039519338,0.0002778444,0.000020838508,0.000009443497,0.000013588328,0.00017803998,0.00006919365,0.00028386476,0.9971988,0.001495452],"study_design_scores_gemma":[0.00050675916,0.000041322808,0.0025580332,0.00020444872,0.00003666151,0.00007035215,0.00009237559,0.0014320344,0.00087870524,0.0024228385,0.9917149,0.00004158539],"about_ca_topic_score_codex":0.013815861,"about_ca_topic_score_gemma":0.02420468,"teacher_disagreement_score":0.119577914,"about_ca_system_score_codex":0.0013953461,"about_ca_system_score_gemma":0.001960089,"threshold_uncertainty_score":0.400028},"labels":[],"label_agreement":null},{"id":"W4398757454","doi":"10.1162/tacl_a_00667","title":"Evaluating Correctness and Faithfulness of Instruction-Following Models for Question Answering","year":2024,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; Mila - Quebec Artificial Intelligence Institute; Minnow Environmental (Canada); Research Canada","funders":"","keywords":"Correctness; Computer science; Question answering; Natural language processing; Artificial intelligence; Programming language","score_opus":0.037789028206600114,"score_gpt":0.3285510699647791,"score_spread":0.290762041758179,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398757454","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6171094,0.001053758,0.3619538,0.001319619,0.00019162724,0.00066373753,0.0010974233,0.012134383,0.004476221],"genre_scores_gemma":[0.9459374,0.00008363485,0.051544532,0.00017416294,0.000027124594,0.00015640388,0.00108523,0.00031009995,0.0006813122],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98790956,0.007259812,0.0008078537,0.0021618367,0.001418576,0.00044220377],"domain_scores_gemma":[0.8792539,0.09748355,0.0043364423,0.01080103,0.0064685987,0.0016565819],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01909782,0.0011155009,0.0009150107,0.0017524393,0.0006676409,0.002926078,0.0023026743,0.002493935,0.002085758],"category_scores_gemma":[0.10400178,0.0005551682,0.0008927653,0.00090172986,0.0016451441,0.0034264047,0.0024253388,0.0025984256,0.0009658962],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044874405,0.0011657058,0.06192508,0.0012092706,0.0006914596,0.00030912,0.0046416195,0.51508117,0.027958756,0.009305319,0.011284415,0.36194074],"study_design_scores_gemma":[0.00004629346,0.00029813134,0.0029318405,0.000042507727,0.00005636044,0.00006380582,0.0001859486,0.98193467,0.009088089,0.0044881315,0.00083025923,0.000033981898],"about_ca_topic_score_codex":0.008827999,"about_ca_topic_score_gemma":0.008541726,"teacher_disagreement_score":0.9809022,"about_ca_system_score_codex":0.0020612844,"about_ca_system_score_gemma":0.0016766933,"threshold_uncertainty_score":0.10100013},"labels":[],"label_agreement":null},{"id":"W4398858936","doi":"10.7910/dvn/u1jec0/uyaoab","title":"DonnellyBJPSreplication.R","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Replication (statistics); Identity (music); Genealogy; Internet privacy; Political science; Biology; Art; History; Computer science; Aesthetics; Virology","score_opus":0.026623134975943548,"score_gpt":0.2424498879123248,"score_spread":0.21582675293638126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398858936","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0001222295,0.00011439586,0.00025094222,0.00013284067,0.00008836219,0.000023849932,0.9957035,0.0023289565,0.0012349033],"genre_scores_gemma":[0.00059830636,0.000103558,0.000973449,0.00013453363,0.000031733187,0.00023057727,0.9960968,0.0007459191,0.0010851063],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972631,0.000717221,0.00026699112,0.0008217014,0.0005253966,0.00040565734],"domain_scores_gemma":[0.99298286,0.0023005907,0.00047248838,0.0025677078,0.0010229222,0.0006532868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003753086,0.0043042633,0.0025336745,0.005707777,0.0016493009,0.0055973013,0.0061423145,0.003247576,0.17657454],"category_scores_gemma":[0.020812593,0.0013631767,0.0026165035,0.007412575,0.0011156172,0.0025046708,0.004549397,0.0029733868,0.25734514],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031811556,0.000011445119,0.00022298911,0.00043115276,0.000028414033,0.000009913363,0.000015346679,0.000101109654,0.00004146438,0.0003669892,0.997547,0.0011924279],"study_design_scores_gemma":[0.00043148157,0.000024241795,0.0013954369,0.0003710169,0.00005472808,0.00006499269,0.00005533946,0.0006433172,0.00040874927,0.003282783,0.9932259,0.000042015614],"about_ca_topic_score_codex":0.015424017,"about_ca_topic_score_gemma":0.026636945,"teacher_disagreement_score":0.17657454,"about_ca_system_score_codex":0.0016359448,"about_ca_system_score_gemma":0.003799385,"threshold_uncertainty_score":0.5907007},"labels":[],"label_agreement":null},{"id":"W4399013580","doi":"10.2196/59680","title":"Is Boundary Annotation Necessary? Evaluating Boundary-Free Approaches to Improve Clinical Named Entity Annotation Efficiency: Case Study","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Annotation; Computer science; Boundary (topology); Information retrieval; Natural language processing; Data mining; Artificial intelligence; Mathematics","score_opus":0.132606346198434,"score_gpt":0.4013234067763562,"score_spread":0.2687170605779222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399013580","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7939178,0.0048371386,0.19173455,0.001156428,0.00024046593,0.0010218574,0.0010253184,0.0025112508,0.0035551612],"genre_scores_gemma":[0.7685136,0.0010712264,0.22601703,0.0002762758,0.000105795945,0.0004641435,0.0022457007,0.00038240163,0.00092383276],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98631525,0.00832463,0.0011540945,0.002307116,0.0015507853,0.00034816898],"domain_scores_gemma":[0.93517923,0.04957693,0.0028036756,0.005261,0.0062356866,0.0009436186],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.022635957,0.001501308,0.0010908867,0.0022787752,0.001280879,0.0021245012,0.0021655818,0.0025771619,0.0012357326],"category_scores_gemma":[0.05578799,0.0005116272,0.0012211122,0.0019304575,0.0012446803,0.0032039923,0.0026871506,0.0014601077,0.00074974395],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0052275043,0.0022209105,0.092435524,0.005149535,0.00088197767,0.0039637554,0.008674208,0.1516732,0.03045985,0.0048238877,0.013359993,0.6811297],"study_design_scores_gemma":[0.0005375718,0.0036385239,0.05118015,0.0008721758,0.0011096718,0.0041845758,0.0051491875,0.812237,0.08915498,0.009220521,0.022362702,0.0003529157],"about_ca_topic_score_codex":0.004567425,"about_ca_topic_score_gemma":0.005239975,"teacher_disagreement_score":0.97736406,"about_ca_system_score_codex":0.0020312509,"about_ca_system_score_gemma":0.0015966538,"threshold_uncertainty_score":0.11971176},"labels":[],"label_agreement":null},{"id":"W4399174990","doi":"10.1145/3654984","title":"Unstructured Data Fusion for Schema and Data Extraction","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ACM on Management of Data","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Information retrieval; Schema (genetic algorithms); Component (thermodynamics); Data mining; Table (database)","score_opus":0.13553878951404624,"score_gpt":0.3611646660040275,"score_spread":0.22562587648998128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399174990","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005283117,0.0010136061,0.97902346,0.00042911866,0.00009096947,0.00027918836,0.0038240163,0.008567123,0.0014894741],"genre_scores_gemma":[0.04773195,0.0006501893,0.9341906,0.00028891323,0.000049265567,0.0002496573,0.0151884165,0.0004975746,0.001153388],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99507093,0.0010032109,0.00054775045,0.0011237198,0.0020647733,0.00018965312],"domain_scores_gemma":[0.98945016,0.0033476572,0.00068684353,0.004246322,0.0020654334,0.00020354305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048620133,0.0015933871,0.001557241,0.0061145136,0.0010770715,0.002766668,0.0025415837,0.001382983,0.0041517504],"category_scores_gemma":[0.018212255,0.00083058147,0.0027428796,0.0076521495,0.0011606152,0.007479253,0.004695519,0.0027574361,0.0036031613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004036581,0.00027856618,0.0057885326,0.0015374925,0.00034487678,0.0008471529,0.0011083203,0.045446634,0.029075319,0.05368097,0.035161097,0.8263273],"study_design_scores_gemma":[0.00008141572,0.00022921874,0.0024565202,0.00029590103,0.00019051133,0.0014471185,0.00090682425,0.62741417,0.08410067,0.14040375,0.14231431,0.00015960379],"about_ca_topic_score_codex":0.0035258203,"about_ca_topic_score_gemma":0.005729317,"teacher_disagreement_score":0.0061145136,"about_ca_system_score_codex":0.0015541754,"about_ca_system_score_gemma":0.0026816295,"threshold_uncertainty_score":0.025713086},"labels":[],"label_agreement":null},{"id":"W4399204106","doi":"10.1007/978-3-031-55642-5_1","title":"An Overview on Large Language Models","year":2024,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Linguistics; Philosophy","score_opus":0.06889453219493441,"score_gpt":0.3030592937623716,"score_spread":0.2341647615674372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399204106","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010260993,0.025187852,0.92646676,0.0024568127,0.00068140257,0.000098152144,0.0026363307,0.008555971,0.03289073],"genre_scores_gemma":[0.05678734,0.063798316,0.71460396,0.002633204,0.0032447795,0.001036254,0.022990728,0.0083233,0.12658209],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991474,0.00028563256,0.000055583507,0.00015379653,0.0003165173,0.00004108516],"domain_scores_gemma":[0.9971596,0.0020873789,0.000054937234,0.00032213848,0.00031751493,0.000058512513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018319853,0.00166029,0.0013626628,0.0021973315,0.00072484306,0.0031248054,0.0016961609,0.0011992331,0.028567817],"category_scores_gemma":[0.005593552,0.0012364379,0.0016921213,0.0044482043,0.0005984871,0.0068050628,0.0015132142,0.0029582868,0.022772538],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005564725,0.000088050656,0.00036654426,0.0008286694,0.000091888316,0.00014245968,0.00019097659,0.023863727,0.0022216742,0.12519103,0.19702546,0.6499338],"study_design_scores_gemma":[0.000017971264,0.000034021312,0.00047557792,0.00027327653,0.00008355856,0.0004642835,0.000087160195,0.15995461,0.0029613066,0.36754447,0.4680377,0.00006607581],"about_ca_topic_score_codex":0.003923901,"about_ca_topic_score_gemma":0.005531864,"teacher_disagreement_score":0.028567817,"about_ca_system_score_codex":0.0011947904,"about_ca_system_score_gemma":0.0015667761,"threshold_uncertainty_score":0.095568895},"labels":[],"label_agreement":null},{"id":"W4399205224","doi":"10.1609/icwsm.v18i1.31324","title":"Reliability Analysis of Psychological Concept Extraction and Classification in User-Penned Text","year":2024,"lang":"en","type":"article","venue":"Proceedings of the International AAAI Conference on Web and Social Media","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute","funders":"National Institute on Aging; National Institutes of Health","keywords":"Witness; Task (project management); Reliability (semiconductor); Perception; Focus (optics); Computer science; Cognitive psychology; Social media; Psychology; Binary classification; Deep learning; Artificial intelligence; Natural language processing; Data science; World Wide Web","score_opus":0.06272372588259084,"score_gpt":0.33273245222357256,"score_spread":0.2700087263409817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399205224","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.86956984,0.0008641017,0.10960998,0.001070361,0.00019421615,0.00038413235,0.010878385,0.0028476275,0.004581323],"genre_scores_gemma":[0.94720465,0.0001364716,0.03931013,0.000097796285,0.00013714352,0.00036622054,0.01145539,0.00015442447,0.0011376963],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99337775,0.002814072,0.00078262575,0.0011638014,0.0016093254,0.00025248138],"domain_scores_gemma":[0.871254,0.10721577,0.006013283,0.0062215254,0.00845799,0.0008374373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009179532,0.0006986194,0.00047564117,0.0051835445,0.00071541104,0.001941134,0.00088246213,0.0011415853,0.002730541],"category_scores_gemma":[0.09201656,0.00024970894,0.00075921504,0.00271893,0.0009337461,0.003104231,0.0020086218,0.0016120817,0.0016291917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028983848,0.00089029846,0.42678654,0.0027850936,0.000539618,0.0013642529,0.018646628,0.019602979,0.029518617,0.010272696,0.028518012,0.45817697],"study_design_scores_gemma":[0.00012457278,0.0006452183,0.35259736,0.0002986588,0.00026457608,0.0010528261,0.005540395,0.564326,0.029732391,0.018975489,0.026222453,0.00021993741],"about_ca_topic_score_codex":0.00219723,"about_ca_topic_score_gemma":0.002233544,"teacher_disagreement_score":0.009179532,"about_ca_system_score_codex":0.0007811617,"about_ca_system_score_gemma":0.0005751084,"threshold_uncertainty_score":0.048546612},"labels":[],"label_agreement":null},{"id":"W4399210979","doi":"10.1007/978-3-031-63028-6_6","title":"Preliminary Systematic Review of Open-Source Large Language Models in Education","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University; Simon Fraser University; Mount Saint Vincent University","funders":"","keywords":"Computer science; Open source; Programming language; Natural language processing; Artificial intelligence; Software engineering; Software","score_opus":0.020518801951125328,"score_gpt":0.28580495849316945,"score_spread":0.2652861565420441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399210979","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007965246,0.9605225,0.013477551,0.0052321996,0.0005862566,0.0020715569,0.0076793656,0.00039533997,0.002069967],"genre_scores_gemma":[0.12748449,0.8062744,0.039587375,0.007510892,0.00047709324,0.00782313,0.009325153,0.00034742005,0.00117001],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9532169,0.029002469,0.009648959,0.0023968895,0.005308664,0.00042613578],"domain_scores_gemma":[0.65867597,0.29963478,0.016701045,0.009185742,0.014450852,0.0013515835],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.065157704,0.0013810648,0.0042889677,0.009919971,0.00075872213,0.0041137994,0.003217247,0.0016772338,0.006242803],"category_scores_gemma":[0.25043842,0.0012475384,0.0073107537,0.008323619,0.0015637021,0.0055001015,0.004338572,0.002568414,0.0010033236],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010519264,0.00019848261,0.004685705,0.57558507,0.01773513,0.000109932436,0.0021492804,0.0011518046,0.0005850677,0.0024841912,0.015882723,0.37838075],"study_design_scores_gemma":[0.000870386,0.0006294632,0.011013293,0.7787098,0.092606895,0.00021486638,0.0011746247,0.0011289405,0.00089853466,0.006567883,0.10602805,0.00015728177],"about_ca_topic_score_codex":0.00885132,"about_ca_topic_score_gemma":0.044629883,"teacher_disagreement_score":0.065157704,"about_ca_system_score_codex":0.0040783123,"about_ca_system_score_gemma":0.023288101,"threshold_uncertainty_score":0.34459102},"labels":[],"label_agreement":null},{"id":"W4399327517","doi":"10.1101/2024.06.03.24308405","title":"Evaluating the Efficacy of Large Language Models for Systematic Review and Meta-Analysis Screening","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Meta-analysis; Systematic review; Computer science; Natural language processing; Psychology; Medicine; MEDLINE; Political science; Internal medicine","score_opus":0.1946645548582948,"score_gpt":0.412796460961748,"score_spread":0.21813190610345318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399327517","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032245863,0.0107131135,0.8461773,0.013671274,0.001415325,0.055250216,0.014044145,0.022048255,0.004434463],"genre_scores_gemma":[0.08113744,0.0011122094,0.86451644,0.0012046753,0.00015443006,0.048886698,0.0019841273,0.00066575577,0.00033828075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.46744525,0.45466623,0.0477652,0.013123005,0.01578971,0.0012106806],"domain_scores_gemma":[0.08626276,0.84675235,0.023387073,0.027708875,0.014936262,0.00095272076],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.5114839,0.004482095,0.0067836223,0.012640317,0.0024841523,0.009209514,0.005160984,0.0045769066,0.012766824],"category_scores_gemma":[0.8218209,0.0038153003,0.020164195,0.013115178,0.0029580195,0.008915904,0.0083823,0.0051696044,0.0026071675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.016364293,0.00092364266,0.025456337,0.18634795,0.039828684,0.001334248,0.006382078,0.10001949,0.003927531,0.031519365,0.053338833,0.5345576],"study_design_scores_gemma":[0.01823894,0.0045702923,0.015756559,0.047284428,0.06024392,0.0012925372,0.0012867794,0.61680716,0.010013696,0.15830712,0.0644216,0.0017770482],"about_ca_topic_score_codex":0.0033416268,"about_ca_topic_score_gemma":0.009836741,"teacher_disagreement_score":0.4885161,"about_ca_system_score_codex":0.005665074,"about_ca_system_score_gemma":0.026274387,"threshold_uncertainty_score":0.60242736},"labels":[],"label_agreement":null},{"id":"W4399426318","doi":"10.1162/tacl_a_00669","title":"Source-Free Domain Adaptation for Question Answering with Masked Self-training","year":2024,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Domain adaptation; Adaptation (eye); Domain (mathematical analysis); Question answering; Training (meteorology); Training set; Natural language processing; Artificial intelligence; Information retrieval; Speech recognition; Psychology","score_opus":0.019462492397112383,"score_gpt":0.25327610768223807,"score_spread":0.23381361528512568,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399426318","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06854468,0.00094272086,0.92134887,0.00034809194,0.00012671965,0.00014781157,0.00029625747,0.006584733,0.0016601612],"genre_scores_gemma":[0.75128883,0.0002728077,0.24208963,0.00076922163,0.000119505785,0.00032625496,0.0016598208,0.00042713893,0.003046776],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987387,0.0005694845,0.00005777819,0.00043018835,0.00012888745,0.00007501158],"domain_scores_gemma":[0.9966581,0.0017167389,0.00015788988,0.000950499,0.00039749802,0.00011918245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026995838,0.0010523901,0.00081868446,0.00076986244,0.00045456633,0.0007461662,0.0019479749,0.0013915416,0.0018513402],"category_scores_gemma":[0.0061693806,0.0005001897,0.0009997393,0.00066329737,0.00089090434,0.0023782344,0.002549531,0.002318628,0.0012718829],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077236263,0.0007603513,0.0064165476,0.00038561365,0.00032110023,0.00026280395,0.0009438813,0.26889172,0.05941187,0.007608148,0.013298191,0.6409274],"study_design_scores_gemma":[0.00001800668,0.000070863265,0.0006737737,0.000013360108,0.000022347123,0.00006229043,0.000045061523,0.9837277,0.008378778,0.0051131817,0.0018596542,0.000014966665],"about_ca_topic_score_codex":0.002303574,"about_ca_topic_score_gemma":0.0025671916,"teacher_disagreement_score":0.0026995838,"about_ca_system_score_codex":0.0006379754,"about_ca_system_score_gemma":0.00071498705,"threshold_uncertainty_score":0.014276922},"labels":[],"label_agreement":null},{"id":"W4399434480","doi":"10.1007/s10639-024-12771-3","title":"Exploring quality criteria and evaluation methods in automated question generation: A comprehensive survey","year":2024,"lang":"en","type":"article","venue":"Education and Information Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Quality (philosophy); Computer science; Educational technology; Evaluation methods; Management science; Data science; Mathematics education; Psychology; Engineering; Reliability engineering; Epistemology","score_opus":0.3294129269693888,"score_gpt":0.48108011046805094,"score_spread":0.15166718349866215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399434480","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15633242,0.38829577,0.4190487,0.010825902,0.00038122167,0.0023747776,0.0022064566,0.0019328404,0.018601937],"genre_scores_gemma":[0.4796125,0.13350849,0.37520987,0.0020448929,0.0005584013,0.0020164785,0.004002997,0.0010405794,0.002005798],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.86612844,0.07727612,0.014491274,0.006341276,0.03414722,0.001615591],"domain_scores_gemma":[0.2530116,0.6787577,0.015979351,0.010848848,0.03972883,0.0016736471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.15849812,0.0017173289,0.0028788887,0.019117668,0.0012845694,0.008432834,0.0041525527,0.002756027,0.0030040487],"category_scores_gemma":[0.3652562,0.0011333773,0.002020766,0.013737546,0.0028011068,0.012780455,0.003513431,0.0022541862,0.0008952125],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039656722,0.0005416819,0.03648796,0.010900146,0.00041131405,0.000037785587,0.0038805131,0.0030766085,0.0019725363,0.007818716,0.004448436,0.9300279],"study_design_scores_gemma":[0.0007424949,0.0049313754,0.23949192,0.07300104,0.004261856,0.002417433,0.033976633,0.16428079,0.042410567,0.09791141,0.33543253,0.0011419135],"about_ca_topic_score_codex":0.0044336687,"about_ca_topic_score_gemma":0.0040989732,"teacher_disagreement_score":0.15849812,"about_ca_system_score_codex":0.0038816936,"about_ca_system_score_gemma":0.006207896,"threshold_uncertainty_score":0.8382282},"labels":[],"label_agreement":null},{"id":"W4399441286","doi":"10.1111/coin.12656","title":"Utilizing passage‐level relevance and kernel pooling for enhancing BERT‐based document reranking","year":2024,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China","keywords":"Computer science; Pooling; Relevance (law); Information retrieval; Ranking (information retrieval); Learning to rank; Artificial intelligence; Kernel (algebra); Rank (graph theory); Machine learning; Data mining; Natural language processing; Mathematics","score_opus":0.06767963765381896,"score_gpt":0.33028697392846934,"score_spread":0.2626073362746504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399441286","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09239567,0.00097371987,0.8996962,0.00018744354,0.00011001776,0.00009128111,0.00020154857,0.004516647,0.0018274678],"genre_scores_gemma":[0.84032446,0.00043130163,0.15154618,0.00013980511,0.00017282454,0.00007776956,0.0006679365,0.00026327002,0.006376359],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993819,0.00017869403,0.000042567917,0.000110337096,0.00019236948,0.00009416711],"domain_scores_gemma":[0.9988907,0.00039034107,0.00013221947,0.00017079803,0.0003268739,0.00008917557],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012878997,0.0008871872,0.0010151414,0.0014642007,0.00033871067,0.00087337225,0.0010460265,0.0006378654,0.0018199699],"category_scores_gemma":[0.0033177133,0.00025625917,0.0005420496,0.0011996152,0.0005043814,0.0021833668,0.0008571461,0.0007453447,0.0010362515],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007007291,0.00039668594,0.0029222353,0.0002392061,0.00014155028,0.00028987642,0.00020880351,0.17242913,0.06855201,0.012672651,0.009170075,0.732277],"study_design_scores_gemma":[0.00002240837,0.00013162658,0.0006177371,0.000005230072,0.000033957014,0.00007322981,0.000020233274,0.97991014,0.014888303,0.003085983,0.0011889017,0.000022255244],"about_ca_topic_score_codex":0.0064291772,"about_ca_topic_score_gemma":0.008845988,"teacher_disagreement_score":0.0064291772,"about_ca_system_score_codex":0.00067478674,"about_ca_system_score_gemma":0.0010301645,"threshold_uncertainty_score":0.012783468},"labels":[],"label_agreement":null},{"id":"W4399447660","doi":"10.48550/arxiv.2406.02969","title":"Filtered not Mixed: Stochastic Filtering-Based Online Gating for Mixture of Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Vector Institute; McMaster University; Natural Sciences and Engineering Research Council of Canada; Government of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Gating; Language model; Artificial intelligence; Psychology; Neuroscience","score_opus":0.08123969319778795,"score_gpt":0.2250849381751377,"score_spread":0.14384524497734974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399447660","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004572059,0.00012509878,0.9937412,0.00011141308,0.000024506036,0.000027363789,0.000047797945,0.00081825117,0.00053233217],"genre_scores_gemma":[0.39725643,0.00034985456,0.5965911,0.00063881965,0.00017898457,0.00028632354,0.000598667,0.00044155738,0.0036581843],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984927,0.00049340917,0.00007464682,0.0003388911,0.0003767454,0.00022360575],"domain_scores_gemma":[0.9969919,0.0020297484,0.00018637822,0.00038879324,0.0002520666,0.00015110637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036431213,0.0014302343,0.0016984983,0.00095786597,0.00074499706,0.0016499638,0.003013783,0.0018186215,0.0032693895],"category_scores_gemma":[0.009344738,0.0009782849,0.0012484072,0.0010490962,0.0010171589,0.0035632236,0.002860575,0.0025655585,0.0009996878],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038056285,0.0002534869,0.0024607775,0.00010375582,0.0001719998,0.00019241105,0.00024854645,0.58893317,0.007694002,0.0631213,0.0043912246,0.3320488],"study_design_scores_gemma":[0.0000094635725,0.00002256775,0.00008508067,0.0000055899973,0.000009355853,0.00001689663,0.0000054355814,0.98683274,0.0009732204,0.0113194985,0.0007101929,0.000010099046],"about_ca_topic_score_codex":0.007052794,"about_ca_topic_score_gemma":0.010084487,"teacher_disagreement_score":0.007052794,"about_ca_system_score_codex":0.0011977821,"about_ca_system_score_gemma":0.002073117,"threshold_uncertainty_score":0.019266903},"labels":[],"label_agreement":null},{"id":"W4399453597","doi":"10.48550/arxiv.2406.04220","title":"BEADs: Bias Evaluation Across Domains","year":2024,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Government of Canada; Canadian Institute for Advanced Research","keywords":"Environmental science","score_opus":0.18843058381648173,"score_gpt":0.3731990703844707,"score_spread":0.18476848656798897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399453597","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24796012,0.017927181,0.5927596,0.006477769,0.002685995,0.0014885375,0.039686166,0.06756082,0.023453845],"genre_scores_gemma":[0.65017533,0.0020682388,0.23114371,0.0029904924,0.0008340175,0.0013236596,0.0964609,0.006222159,0.008781351],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9778901,0.013311571,0.0013674771,0.003828055,0.0028775833,0.0007251726],"domain_scores_gemma":[0.96240467,0.023130868,0.001260385,0.008102426,0.0041114707,0.0009900954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02552653,0.0034642909,0.0020391888,0.0038347545,0.001906915,0.003950456,0.003277662,0.0036093758,0.006041033],"category_scores_gemma":[0.063535154,0.0007728427,0.0022236693,0.002776356,0.0019509734,0.0049233674,0.005542905,0.003912132,0.0061609456],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038405214,0.0010494087,0.05721708,0.0024347366,0.0015847206,0.0005175964,0.001134269,0.1726852,0.007787799,0.01785543,0.26334,0.47055328],"study_design_scores_gemma":[0.00088089466,0.000713259,0.0104796905,0.0006007827,0.00048841006,0.00064857095,0.00066522596,0.84258825,0.022174474,0.05707743,0.06347428,0.00020875207],"about_ca_topic_score_codex":0.008421136,"about_ca_topic_score_gemma":0.01006288,"teacher_disagreement_score":0.02552653,"about_ca_system_score_codex":0.002489487,"about_ca_system_score_gemma":0.0031129124,"threshold_uncertainty_score":0.1349988},"labels":[],"label_agreement":null},{"id":"W4399732994","doi":"10.22148/001c.116368","title":"Exploring Gender Differences in Fatwa through Machine Learning","year":2024,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Popularity; Context (archaeology); Computer science; Preprocessor; Margin (machine learning); Artificial intelligence; Machine learning; Thematic analysis; Data science; Psychology; Social psychology; Qualitative research; Sociology; Social science","score_opus":0.2739161805313298,"score_gpt":0.31646283693417326,"score_spread":0.04254665640284344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399732994","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94234127,0.0023904752,0.020688627,0.0011400435,0.00018921478,0.0002336932,0.022804156,0.0005974342,0.009615019],"genre_scores_gemma":[0.9602129,0.00051304785,0.01328368,0.00022573814,0.00013231754,0.00023686749,0.021168094,0.00008887548,0.0041384916],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99849606,0.00053323014,0.00013402852,0.0003975513,0.0002540605,0.00018508095],"domain_scores_gemma":[0.9865471,0.009588533,0.0013281411,0.00084130594,0.0012829839,0.000411944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029431598,0.00054861885,0.0004546012,0.002912832,0.0006712321,0.0017313928,0.00056087406,0.00066314026,0.0042529306],"category_scores_gemma":[0.016294299,0.0001611283,0.00059264334,0.0021807088,0.000398989,0.002227243,0.0009250794,0.0008787919,0.0026500216],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010270717,0.00047188942,0.8156814,0.00070904085,0.00016851173,0.00034644853,0.005629693,0.002009499,0.009153786,0.0031856203,0.016370613,0.14524634],"study_design_scores_gemma":[0.000060022907,0.00049448473,0.84236556,0.000341706,0.00017298182,0.0010809746,0.013995109,0.065797,0.007772036,0.0106280185,0.05719349,0.00009869423],"about_ca_topic_score_codex":0.0049723536,"about_ca_topic_score_gemma":0.009101549,"teacher_disagreement_score":0.0049723536,"about_ca_system_score_codex":0.0006045235,"about_ca_system_score_gemma":0.00049401075,"threshold_uncertainty_score":0.015565097},"labels":[],"label_agreement":null},{"id":"W4399768519","doi":"10.1109/tse.2024.3411928","title":"LUNA: A Model-Based Universal Analysis Framework for Large Language Models","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Japan Society for the Promotion of Science; JST-Mirai Program; Alberta Machine Intelligence Institute; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Programming language; Modeling language; Software engineering; Data science; Natural language processing; Software","score_opus":0.013074621913195335,"score_gpt":0.2445529802957176,"score_spread":0.23147835838252226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399768519","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00039830257,0.00013828327,0.9952748,0.0001397854,0.000019545334,0.00004092865,0.00023599039,0.0030244747,0.00072789704],"genre_scores_gemma":[0.06763932,0.0006280428,0.9226523,0.00039112303,0.00016772332,0.0006191774,0.0019837357,0.0023873148,0.0035311794],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99726164,0.001153142,0.0002281109,0.00044041077,0.0007125275,0.00020411565],"domain_scores_gemma":[0.99730885,0.001543475,0.00024575205,0.00042709944,0.00035655405,0.000118271135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003891363,0.0017338862,0.0013570401,0.002961291,0.0012126208,0.0035921584,0.0038001502,0.0014179994,0.008856321],"category_scores_gemma":[0.0098814415,0.0015350868,0.0049431426,0.0015380285,0.0018064566,0.0054466887,0.004421318,0.0040963735,0.0036351888],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012202023,0.00007548814,0.001037729,0.00038227357,0.0003087228,0.00040863242,0.0005345108,0.1461991,0.0025208506,0.7089982,0.016222944,0.1231895],"study_design_scores_gemma":[0.000015764632,0.000020074725,0.00008824165,0.00004928518,0.00003949091,0.00006975132,0.000040444997,0.73667455,0.0007987951,0.24363953,0.018539276,0.000024743154],"about_ca_topic_score_codex":0.0098955175,"about_ca_topic_score_gemma":0.01329402,"teacher_disagreement_score":0.0098955175,"about_ca_system_score_codex":0.0019658005,"about_ca_system_score_gemma":0.002833168,"threshold_uncertainty_score":0.029627383},"labels":[],"label_agreement":null},{"id":"W4399802498","doi":"10.1016/j.asoc.2024.111893","title":"Class conditioned text generation with style attention mechanism for embracing diversity","year":2024,"lang":"en","type":"article","venue":"Applied Soft Computing","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; National Research Foundation of Korea; Information Technology Research Centre; Ministry of Science, ICT and Future Planning","keywords":"Computer science; Fluency; Natural language generation; Artificial intelligence; Transformer; Class (philosophy); Adversarial system; Scope (computer science); Text generation; Natural language processing; Natural language; Language model; Machine learning; Linguistics; Programming language","score_opus":0.023586817822720024,"score_gpt":0.2337100973111561,"score_spread":0.2101232794884361,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399802498","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06683505,0.0005375723,0.91982883,0.00044279924,0.0004317892,0.0002601613,0.00035796754,0.006536732,0.0047690035],"genre_scores_gemma":[0.7570473,0.00021642829,0.23193878,0.00034061322,0.0004128332,0.0002706546,0.00080215634,0.0006735414,0.0082976995],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914944,0.00020304225,0.00004578803,0.00031667555,0.00020100406,0.00008401108],"domain_scores_gemma":[0.9978364,0.00088216556,0.00009389873,0.000374758,0.000631197,0.000181683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001198766,0.0006934633,0.00068310025,0.0009974854,0.0005876555,0.000938166,0.0013236678,0.0009623211,0.005488124],"category_scores_gemma":[0.0046318127,0.00028595462,0.00081297685,0.00083874,0.0003904944,0.0014908783,0.0013546803,0.0013820102,0.00148433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085656677,0.0005085911,0.004089196,0.00024046445,0.00013432902,0.00023616207,0.00046911664,0.027073333,0.1123473,0.012740724,0.016003823,0.8253003],"study_design_scores_gemma":[0.00010611827,0.00023324312,0.002848184,0.000018727578,0.00012623121,0.00019869454,0.00007492674,0.93534106,0.041390024,0.012981051,0.00663195,0.000049790833],"about_ca_topic_score_codex":0.001630881,"about_ca_topic_score_gemma":0.0023135582,"teacher_disagreement_score":0.005488124,"about_ca_system_score_codex":0.00043242986,"about_ca_system_score_gemma":0.0008501318,"threshold_uncertainty_score":0.018359601},"labels":[],"label_agreement":null},{"id":"W4400025652","doi":"10.17615/zfj3-7992","title":"Case study in major quotation errors: a critical commentary on the Newcastle–Ottawa scale","year":2024,"lang":"en","type":"article","venue":"UNC Libraries","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Federalno Ministarstvo Obrazovanja i Nauke; Bundesministerium für Bildung und Forschung","keywords":"Scale (ratio); History; Computer science; Geography; Cartography","score_opus":0.05108569341887,"score_gpt":0.2957474867031689,"score_spread":0.2446617932842989,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400025652","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017517579,0.0021864423,0.00068518525,0.97280145,0.017571675,0.00004992036,0.000036924423,0.000027975344,0.0048887916],"genre_scores_gemma":[0.06579158,0.003671378,0.003499326,0.8822693,0.020353712,0.00036494728,0.000050766794,0.0004623132,0.02353662],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.8744216,0.051467117,0.018133562,0.012380362,0.03386847,0.0097289095],"domain_scores_gemma":[0.48248446,0.3917452,0.015317293,0.009928327,0.088070296,0.012454344],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08969706,0.0011676552,0.0021091448,0.0067453603,0.042343613,0.026015542,0.012402303,0.055397358,0.007650177],"category_scores_gemma":[0.3416992,0.0024082381,0.0020694959,0.0069137746,0.03968625,0.020805443,0.015376288,0.067958735,0.0022992406],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000052062347,0.00001636138,0.00042149136,0.0008000377,0.0000329303,0.0023109314,0.11016171,0.00011357568,0.0003395154,0.039425064,0.83779186,0.008534504],"study_design_scores_gemma":[0.0000153621,0.000012626834,0.0010555339,0.003484774,0.000049865903,0.0004861551,0.110902674,0.00018352664,0.00039854867,0.007451426,0.8757694,0.00019014881],"about_ca_topic_score_codex":0.30202898,"about_ca_topic_score_gemma":0.49744707,"teacher_disagreement_score":0.91030294,"about_ca_system_score_codex":0.060983647,"about_ca_system_score_gemma":0.07846915,"threshold_uncertainty_score":0.60054195},"labels":[],"label_agreement":null},{"id":"W4400046281","doi":"10.1007/s10791-024-09435-8","title":"Injecting the score of the first-stage retriever as text improves BERT-based re-rankers","year":2024,"lang":"en","type":"article","venue":"Discover Computing","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Horizon 2020 Framework Programme","keywords":"Labrador Retriever; Stage (stratigraphy); Computer science; Natural language processing; Artificial intelligence; Information retrieval; Medicine; Biology; Surgery","score_opus":0.01901571512607659,"score_gpt":0.2539155939652864,"score_spread":0.23489987883920982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400046281","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20725876,0.0013048896,0.75552243,0.0003526224,0.00019855345,0.00036337718,0.0007959142,0.026295671,0.007907754],"genre_scores_gemma":[0.6904309,0.00041474742,0.2845199,0.00023362807,0.00012884289,0.00013458471,0.0022489112,0.0011688748,0.020719774],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984712,0.00039377398,0.0001130911,0.00031560435,0.0005308605,0.00017543744],"domain_scores_gemma":[0.9969633,0.0011007406,0.00022631894,0.0008697653,0.0007234916,0.0001162881],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023612517,0.0014219485,0.0011439289,0.0017029486,0.0004319219,0.0014175544,0.001605161,0.00079344853,0.0068882643],"category_scores_gemma":[0.007640854,0.0004549825,0.00080346974,0.00086930575,0.0005113527,0.0029982335,0.0015119093,0.0011839799,0.0054499675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009751199,0.00047754086,0.0048060883,0.0003091051,0.00014882378,0.00020304296,0.000241165,0.0765998,0.05248222,0.0047947373,0.006737557,0.8522247],"study_design_scores_gemma":[0.00011415487,0.000766783,0.0029366633,0.000032508084,0.00016360509,0.0004134461,0.00017548181,0.8707302,0.11313927,0.0035880946,0.007845073,0.00009472928],"about_ca_topic_score_codex":0.00493558,"about_ca_topic_score_gemma":0.008570931,"teacher_disagreement_score":0.0068882643,"about_ca_system_score_codex":0.0006839171,"about_ca_system_score_gemma":0.0012287427,"threshold_uncertainty_score":0.023043573},"labels":[],"label_agreement":null},{"id":"W4400083089","doi":"10.2139/ssrn.4878306","title":"Entity and Relation Extractions for Threat Intelligence Knowledge Graphs","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Relation (database); Knowledge graph; Computer science; Business; Knowledge management; Artificial intelligence; Data mining","score_opus":0.0293684659563249,"score_gpt":0.30660768789122333,"score_spread":0.2772392219348984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400083089","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05238233,0.0027311035,0.86959636,0.0015996487,0.00018551685,0.00073347177,0.04712985,0.013789329,0.011852325],"genre_scores_gemma":[0.33062175,0.0027215146,0.55939424,0.0002830019,0.00024071669,0.00049198314,0.09753862,0.0010557378,0.0076524206],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985935,0.0003107912,0.00017625875,0.0004412133,0.0003553418,0.00012287594],"domain_scores_gemma":[0.9947024,0.0034358175,0.00036801441,0.00084971834,0.0004959009,0.00014814363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013286485,0.00097197056,0.00083929877,0.009543034,0.0012432182,0.0026527857,0.0011851906,0.0014478652,0.007544334],"category_scores_gemma":[0.010574095,0.00066593266,0.0021538106,0.009006437,0.00044099538,0.0050647897,0.0015448027,0.0016587424,0.0037055977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053028465,0.00044199536,0.012562736,0.0016216205,0.0004956212,0.0012983374,0.0015844457,0.03172911,0.014920068,0.063004196,0.06659629,0.8052153],"study_design_scores_gemma":[0.00013177564,0.00019315034,0.021782396,0.0006516851,0.0009926417,0.0024792755,0.0017667203,0.51214373,0.027214179,0.28710258,0.1453893,0.00015254626],"about_ca_topic_score_codex":0.005379437,"about_ca_topic_score_gemma":0.0109246345,"teacher_disagreement_score":0.009543034,"about_ca_system_score_codex":0.00079574774,"about_ca_system_score_gemma":0.001476364,"threshold_uncertainty_score":0.025238335},"labels":[],"label_agreement":null},{"id":"W4400153822","doi":"10.1007/978-3-031-60916-9_9","title":"Using Similarity Based on Embeddings","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in social networks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Similarity (geometry); Computer science; Artificial intelligence","score_opus":0.03938916913820685,"score_gpt":0.2897548964617847,"score_spread":0.25036572732357787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400153822","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03370264,0.0018631212,0.9540151,0.00038804184,0.00042588275,0.00013188692,0.0008375479,0.0027401496,0.0058956114],"genre_scores_gemma":[0.5239207,0.0017038761,0.45002365,0.00028098395,0.0008295638,0.00024412914,0.0067759645,0.0012423489,0.014978803],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99733764,0.0007939037,0.00018976345,0.00073238096,0.0007952252,0.00015104956],"domain_scores_gemma":[0.9951342,0.0021422522,0.0003320756,0.0013746312,0.0008812893,0.00013556698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001489038,0.0013160971,0.0020654034,0.0044054557,0.0006271441,0.002316163,0.001758911,0.0016581346,0.0055828155],"category_scores_gemma":[0.01034629,0.00064555416,0.0014483739,0.004692516,0.00065005466,0.0076129017,0.0030156116,0.0016782425,0.0038974392],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053852197,0.00036075784,0.0035567037,0.0003873522,0.00039820524,0.0001240938,0.0002083872,0.040209103,0.012249108,0.03932313,0.016699767,0.8859448],"study_design_scores_gemma":[0.000046033856,0.00023800344,0.0015202768,0.00007217727,0.00012997795,0.00034733352,0.00016611358,0.87693393,0.007026184,0.10402345,0.009432185,0.00006419947],"about_ca_topic_score_codex":0.0018838223,"about_ca_topic_score_gemma":0.002197135,"teacher_disagreement_score":0.0055828155,"about_ca_system_score_codex":0.0005051854,"about_ca_system_score_gemma":0.0004390347,"threshold_uncertainty_score":0.01867634},"labels":[],"label_agreement":null},{"id":"W4400267753","doi":"10.1145/3649217.3653554","title":"Can Small Language Models With Retrieval-Augmented Generation Replace Large Language Models When Learning Computer Science?","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Universitas Brawijaya","keywords":"Computer science; Language model; Natural language processing; Artificial intelligence; Universal Networking Language; Information retrieval; Natural language; Comprehension approach","score_opus":0.026593168225165588,"score_gpt":0.24822958569939937,"score_spread":0.2216364174742338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400267753","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031931665,0.0021074628,0.9300785,0.010215127,0.0012246671,0.0002613504,0.0008556867,0.012837817,0.010487718],"genre_scores_gemma":[0.46167392,0.0013927573,0.51714903,0.003962363,0.000583708,0.00040142058,0.0023094057,0.003081571,0.0094457725],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99390125,0.0038502389,0.00025663097,0.0009713951,0.00066228356,0.00035822415],"domain_scores_gemma":[0.97439665,0.015852554,0.0006339857,0.0070629893,0.0016582892,0.00039560389],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009797305,0.0012269135,0.0012908698,0.001066319,0.0007806824,0.0034939991,0.0032656547,0.0025785717,0.008013146],"category_scores_gemma":[0.04942959,0.0010802925,0.0018988358,0.0009653681,0.0019030382,0.016153336,0.0027794377,0.0038667305,0.007921291],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017388727,0.0006405884,0.0066071176,0.00088206294,0.0003974969,0.00040789615,0.0020542638,0.08571647,0.010412324,0.091525614,0.04324065,0.75637656],"study_design_scores_gemma":[0.0002687241,0.00029996128,0.00079776667,0.00019890128,0.0001924454,0.00040116205,0.0004557614,0.7286003,0.008508846,0.22125234,0.038887084,0.00013674209],"about_ca_topic_score_codex":0.010235299,"about_ca_topic_score_gemma":0.011868159,"teacher_disagreement_score":0.010235299,"about_ca_system_score_codex":0.0012153572,"about_ca_system_score_gemma":0.0026320405,"threshold_uncertainty_score":0.05181372},"labels":[],"label_agreement":null},{"id":"W4400280941","doi":"10.32473/flairs.37.1.135597","title":"Exploration of Word Embeddings with Graph-Based Context Adaptation for Enhanced Word Vectors","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Word embedding; Word (group theory); Natural language understanding; Embedding; Context (archaeology); Natural language; Graph; Representation (politics); Semantic similarity; Linguistics; Theoretical computer science","score_opus":0.19174769746374895,"score_gpt":0.3730239915487573,"score_spread":0.18127629408500837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400280941","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09261729,0.0010957106,0.899986,0.00029771606,0.00012416122,0.00012008483,0.0004881723,0.0030561327,0.0022147188],"genre_scores_gemma":[0.5304945,0.00069290673,0.46256247,0.0001828326,0.000080509235,0.0002205156,0.0021208155,0.00045551013,0.0031898785],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99978775,0.00006479519,0.000015436495,0.00007968007,0.000032421318,0.000019890422],"domain_scores_gemma":[0.99937314,0.0003366145,0.00006017143,0.00009679256,0.00010702909,0.000026350705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00036540482,0.0010650901,0.00064223364,0.0014350215,0.00028194266,0.0008429093,0.0007122206,0.00057439983,0.002165892],"category_scores_gemma":[0.0024006057,0.00032610513,0.0007790846,0.0017200042,0.00039261417,0.0021183104,0.0011025623,0.0007924492,0.0009015507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029922687,0.00030056006,0.0029767808,0.0004917214,0.00013547704,0.00020989259,0.0005667627,0.14280479,0.038887624,0.017890671,0.009501737,0.78593475],"study_design_scores_gemma":[0.000026705266,0.00013929448,0.0004947939,0.000024232038,0.000043350687,0.00009312941,0.00016390707,0.9685942,0.0053194873,0.020804876,0.004274895,0.000021038768],"about_ca_topic_score_codex":0.002340098,"about_ca_topic_score_gemma":0.004691986,"teacher_disagreement_score":0.002340098,"about_ca_system_score_codex":0.00028026977,"about_ca_system_score_gemma":0.00045909165,"threshold_uncertainty_score":0.0072456},"labels":[],"label_agreement":null},{"id":"W4400281282","doi":"10.32473/flairs.37.1.135043","title":"Latent Beta-Liouville Probabilistic Modeling for Bursty Topic Discovery in Textual Data","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Latent Dirichlet allocation; Burstiness; Perplexity; Computer science; Topic model; Natural language processing; Language model; Word (group theory); Probabilistic logic; Dirichlet distribution; Artificial intelligence; Range (aeronautics); Mathematics","score_opus":0.312632176089768,"score_gpt":0.4053127273422721,"score_spread":0.09268055125250407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400281282","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015000427,0.00072666863,0.9825211,0.00033271505,0.000035416422,0.00007918816,0.00035393747,0.00048506368,0.00046554563],"genre_scores_gemma":[0.5462656,0.0021454461,0.44058505,0.00049264735,0.0004507722,0.0012764129,0.0033426473,0.000364822,0.0050766026],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970432,0.001615011,0.00018055252,0.00064881827,0.0003556696,0.00015685582],"domain_scores_gemma":[0.9882549,0.00969338,0.000696808,0.0006383503,0.0005279176,0.00018868342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062807496,0.0010711291,0.0014542035,0.0031589412,0.0009918342,0.0024577908,0.0027022916,0.0017162487,0.0017718662],"category_scores_gemma":[0.017938854,0.000863705,0.0019153084,0.003793295,0.001209255,0.0043401825,0.0018169503,0.0031206687,0.0012855473],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005610827,0.00030610504,0.012146407,0.0005193325,0.00036153826,0.00039186602,0.002071893,0.61984515,0.007733351,0.12964632,0.0067791273,0.2196378],"study_design_scores_gemma":[0.000009448769,0.000013519281,0.00028104478,0.000012514004,0.000010080364,0.000030440762,0.000030433066,0.96930885,0.00028951958,0.029068304,0.0009325783,0.000013235333],"about_ca_topic_score_codex":0.006063833,"about_ca_topic_score_gemma":0.008995128,"teacher_disagreement_score":0.0062807496,"about_ca_system_score_codex":0.00177307,"about_ca_system_score_gemma":0.0015464099,"threshold_uncertainty_score":0.03321618},"labels":[],"label_agreement":null},{"id":"W4400320688","doi":"10.1007/s10489-024-05430-0","title":"Learning contextual representations for entity retrieval","year":2024,"lang":"en","type":"article","venue":"Applied Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Natural language processing; Artificial intelligence","score_opus":0.04410917281330834,"score_gpt":0.31941305577371903,"score_spread":0.2753038829604107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400320688","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060981918,0.0059861382,0.91951084,0.0013039482,0.00028993253,0.0001738518,0.0029848157,0.0060511357,0.002717379],"genre_scores_gemma":[0.66466486,0.0034437235,0.30903736,0.00050703547,0.0007065121,0.00038020607,0.017123753,0.000562189,0.0035742219],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99855226,0.0004796722,0.00013996928,0.00046760126,0.0001875153,0.00017293857],"domain_scores_gemma":[0.996317,0.002220958,0.00023327873,0.0007168585,0.0003825776,0.00012934678],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018120016,0.0012398809,0.0014744322,0.0037797436,0.0009006702,0.0018442895,0.0019065974,0.001814886,0.003606177],"category_scores_gemma":[0.008960953,0.00073012756,0.0017107103,0.0042292997,0.0005917391,0.005257048,0.0021004295,0.0025965478,0.0018805519],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009781801,0.00073350465,0.00690094,0.0006990198,0.0005218742,0.0003518454,0.0005525332,0.10175331,0.009885447,0.036793277,0.044487126,0.7963429],"study_design_scores_gemma":[0.000083134786,0.00013946324,0.0015565662,0.000100365614,0.00023335504,0.0001490639,0.00020629933,0.9131565,0.003894904,0.07256913,0.00786805,0.000043088443],"about_ca_topic_score_codex":0.0057372274,"about_ca_topic_score_gemma":0.01076732,"teacher_disagreement_score":0.0057372274,"about_ca_system_score_codex":0.0010571213,"about_ca_system_score_gemma":0.0014645834,"threshold_uncertainty_score":0.012063861},"labels":[],"label_agreement":null},{"id":"W4400343221","doi":"10.48550/arxiv.2407.00541","title":"Answering real-world clinical questions using large language model based systems","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"","keywords":"Computer science; Question answering; Natural language processing","score_opus":0.1487119472503954,"score_gpt":0.27807767755163165,"score_spread":0.12936573030123624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400343221","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.053214267,0.002655005,0.9141831,0.0053946543,0.00018679768,0.0018716512,0.003729665,0.01637257,0.0023923286],"genre_scores_gemma":[0.25453752,0.0008539163,0.73555607,0.0015796354,0.00023427853,0.0011936433,0.0047737355,0.00028319086,0.0009879884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9853123,0.010297562,0.0014568464,0.0016031782,0.0011255123,0.00020458201],"domain_scores_gemma":[0.90330034,0.085556544,0.0035138521,0.0038616247,0.0031654886,0.00060205336],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020638103,0.0015691426,0.0012955175,0.0046241027,0.0009593707,0.0045198696,0.002036396,0.002657248,0.0033552037],"category_scores_gemma":[0.064443275,0.0007162907,0.0021581796,0.0023163809,0.0008980848,0.0034918084,0.002938768,0.00181502,0.0015269673],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015980949,0.001121857,0.01701904,0.0049901954,0.0018982345,0.0015029509,0.0056937886,0.17182378,0.02945524,0.015221897,0.027309688,0.7223652],"study_design_scores_gemma":[0.00053552905,0.00057439005,0.0036672351,0.0005608612,0.0007196688,0.0005153119,0.0015333934,0.8826836,0.018192261,0.0615567,0.029185783,0.00027524168],"about_ca_topic_score_codex":0.0035639007,"about_ca_topic_score_gemma":0.0074560796,"teacher_disagreement_score":0.020638103,"about_ca_system_score_codex":0.0020805874,"about_ca_system_score_gemma":0.002806833,"threshold_uncertainty_score":0.10914606},"labels":[],"label_agreement":null},{"id":"W4400348335","doi":"10.32473/flairs.37.1.135561","title":"Abstractive Text Summarization Based on Neural Fusion","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"Natural Sciences and Engineering Research Council of Canada; Alberta Innovates; University of Lethbridge","keywords":"Automatic summarization; Computer science; Natural language processing; Artificial intelligence; Text graph; Graph; Segmentation; Information retrieval; Baseline (sea); Selection (genetic algorithm); Multi-document summarization; Theoretical computer science","score_opus":0.13458505484915448,"score_gpt":0.3708120438842319,"score_spread":0.2362269890350774,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400348335","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04550662,0.0039605103,0.92939895,0.00080963084,0.00033435595,0.00032084828,0.0016051,0.01344642,0.0046176426],"genre_scores_gemma":[0.49276003,0.0023326608,0.47545752,0.0005226022,0.0007127582,0.00051558344,0.009062742,0.0007073538,0.01792875],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994585,0.00009761429,0.000052816333,0.00020913284,0.00013192056,0.00005013194],"domain_scores_gemma":[0.9989196,0.00038266354,0.00017343633,0.00013173456,0.00034285028,0.000049809427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00080862554,0.0016208552,0.0010509489,0.0024270734,0.00050087104,0.001244349,0.0013287719,0.001058272,0.0027713808],"category_scores_gemma":[0.0027928636,0.00034709787,0.001275579,0.0015714944,0.00042972167,0.002407686,0.0010245943,0.0014208739,0.0019944592],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041206184,0.00018788433,0.0010276991,0.000516136,0.00018916876,0.0002580327,0.00041561975,0.10706815,0.042228557,0.0047984743,0.014509774,0.8283884],"study_design_scores_gemma":[0.000043277632,0.00022131631,0.0010570628,0.000045920457,0.00017688613,0.00008928771,0.00011909632,0.95221215,0.023142714,0.011721431,0.011136605,0.00003427966],"about_ca_topic_score_codex":0.0039620725,"about_ca_topic_score_gemma":0.00633918,"teacher_disagreement_score":0.0039620725,"about_ca_system_score_codex":0.00089527684,"about_ca_system_score_gemma":0.0007540545,"threshold_uncertainty_score":0.0092712045},"labels":[],"label_agreement":null},{"id":"W4400399317","doi":"10.1101/2024.07.05.24309412","title":"OQA : A question-answering dataset on orthodontic literature","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Université de Montréal","funders":"","keywords":"Question answering; Information retrieval; Computer science","score_opus":0.030065808405332274,"score_gpt":0.3026635983425187,"score_spread":0.2725977899371864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400399317","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07436412,0.007994188,0.013158184,0.002862426,0.00044560598,0.0011237812,0.88309133,0.009718047,0.0072423727],"genre_scores_gemma":[0.06715475,0.0013133711,0.03272334,0.00057552964,0.00016872576,0.0010590381,0.89303243,0.00021767664,0.0037552144],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985877,0.00038679765,0.00028490226,0.0003521094,0.00029791056,0.00009054066],"domain_scores_gemma":[0.9942893,0.003458095,0.00040959663,0.00061079825,0.00089948217,0.0003326282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020495749,0.0010399489,0.0006388458,0.005668529,0.00074887887,0.0010663976,0.0017890066,0.0024256306,0.011437085],"category_scores_gemma":[0.012361177,0.00028281682,0.0012243073,0.0037216365,0.0005495431,0.0011007657,0.0021771602,0.0010078719,0.0070582926],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012495158,0.0007443439,0.020298243,0.010920897,0.00027832115,0.001241581,0.0012868874,0.00859943,0.015433508,0.0033032787,0.7128329,0.22381105],"study_design_scores_gemma":[0.0009844364,0.00082010776,0.08183191,0.0014323873,0.00028794873,0.002526206,0.001995697,0.0475409,0.014759957,0.007817053,0.8398073,0.00019613835],"about_ca_topic_score_codex":0.009735346,"about_ca_topic_score_gemma":0.017868176,"teacher_disagreement_score":0.011437085,"about_ca_system_score_codex":0.0012170422,"about_ca_system_score_gemma":0.0022165878,"threshold_uncertainty_score":0.038260877},"labels":[],"label_agreement":null},{"id":"W4400406053","doi":"10.2139/ssrn.4856254","title":"Filtered not Mixed: Stochastic Filtering-Based Online Gating for Mixture of Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; Vector Institute; McMaster University; University of Toronto","funders":"","keywords":"Computer science; Gating; Language model; Artificial intelligence; Natural language processing; Psychology; Neuroscience","score_opus":0.021358632563169777,"score_gpt":0.28695927036475916,"score_spread":0.2656006378015894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400406053","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066154595,0.00018467069,0.99116427,0.00017270686,0.00007493224,0.00004208195,0.00016693298,0.0011615205,0.0004173665],"genre_scores_gemma":[0.4147515,0.0007124802,0.5703058,0.00088994566,0.0006520793,0.0005781192,0.0031333643,0.0013266727,0.0076498897],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99788827,0.0009040717,0.000114370254,0.0004499151,0.00034735873,0.00029609306],"domain_scores_gemma":[0.99133885,0.006694763,0.00030727364,0.00079230394,0.0004803546,0.00038643653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049814163,0.001582171,0.0023714944,0.0014037377,0.0009517496,0.002331505,0.0033410457,0.0026478083,0.0059776492],"category_scores_gemma":[0.017633978,0.001691391,0.0016693615,0.0019520778,0.0010379791,0.0044180867,0.0041889837,0.0034472668,0.0021775363],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019408368,0.00062958023,0.002779899,0.00040954014,0.00030853393,0.00040018215,0.00047358705,0.36153346,0.017961616,0.09927269,0.014081117,0.5002089],"study_design_scores_gemma":[0.000030898642,0.00003712132,0.00016196542,0.000009902205,0.00002103432,0.000032547418,0.000008589151,0.97697085,0.0011034238,0.020842636,0.0007653058,0.000015707967],"about_ca_topic_score_codex":0.005479844,"about_ca_topic_score_gemma":0.009750675,"teacher_disagreement_score":0.0059776492,"about_ca_system_score_codex":0.000974547,"about_ca_system_score_gemma":0.0026472404,"threshold_uncertainty_score":0.026344538},"labels":[],"label_agreement":null},{"id":"W4400434219","doi":"10.48550/arxiv.2407.03951","title":"Uncertainty-Guided Likelihood Tree Search","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Bildung und Forschung; International Max Planck Research School for Advanced Methods in Process and Systems Engineering; Deutsche Forschungsgemeinschaft; International Max Planck Research School for Environmental, Cellular and Molecular Microbiology; Government of Canada; Canadian Institute for Advanced Research; Microsoft Research","keywords":"Computer science; Language model; Artificial intelligence","score_opus":0.10853411783174845,"score_gpt":0.21919514072126473,"score_spread":0.11066102288951628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400434219","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013790388,0.00053846696,0.9817499,0.00045552678,0.000036691883,0.0000690402,0.00024075531,0.000852734,0.0022665854],"genre_scores_gemma":[0.43844092,0.0005038867,0.5552651,0.00040920853,0.000098954035,0.00039322797,0.0011943247,0.0005254835,0.0031688705],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985008,0.00066988135,0.00006202315,0.00030812173,0.00032099537,0.00013810961],"domain_scores_gemma":[0.9899776,0.0083763525,0.0004435658,0.0004596917,0.0004965826,0.00024608258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022298906,0.0010181973,0.0018773226,0.0015678186,0.0008895747,0.0014321717,0.0021630586,0.0021406403,0.004784527],"category_scores_gemma":[0.017948214,0.00075077283,0.0011230813,0.002372132,0.0011914229,0.0031618464,0.0022837166,0.002481776,0.0013165],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022251872,0.00008731111,0.002190593,0.0001947785,0.000082914,0.00016496133,0.0003135742,0.78833985,0.0011121003,0.06638491,0.009192594,0.1317139],"study_design_scores_gemma":[0.000024028735,0.000016601241,0.000078480545,0.000012919224,0.000008101292,0.00002938141,0.000016131698,0.9640772,0.00021108145,0.03462287,0.00089556474,0.00000762496],"about_ca_topic_score_codex":0.00561529,"about_ca_topic_score_gemma":0.007683051,"teacher_disagreement_score":0.00561529,"about_ca_system_score_codex":0.0012817688,"about_ca_system_score_gemma":0.0027498507,"threshold_uncertainty_score":0.016005814},"labels":[],"label_agreement":null},{"id":"W4400524635","doi":"10.1145/3626772.3657674","title":"Embark on DenseQuest: A System for Selecting the Best Dense Retriever for a Custom Collection","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Upload; Relevance (law); Ranking (information retrieval); Labrador Retriever; Information retrieval; Selection (genetic algorithm); Data collection; ENCODE; Cloud computing; World Wide Web; Data science; Artificial intelligence; Operating system","score_opus":0.03449104065488361,"score_gpt":0.2781984682090163,"score_spread":0.24370742755413266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400524635","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029058808,0.0014367299,0.24206913,0.001349673,0.00041365818,0.0010922698,0.025561063,0.6713559,0.027662767],"genre_scores_gemma":[0.17260541,0.0012495106,0.6430271,0.0021226816,0.00048400616,0.0011710873,0.09181273,0.04928998,0.038237337],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988458,0.00019734376,0.000080029284,0.00036080956,0.00039796846,0.00011811443],"domain_scores_gemma":[0.9970799,0.0010364671,0.0001354636,0.0008861557,0.00045214396,0.0004097996],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022117882,0.00221845,0.0014545961,0.0021926297,0.0013305234,0.0028743942,0.0022709884,0.0014135459,0.031788997],"category_scores_gemma":[0.007070564,0.0010628753,0.001462305,0.0016639084,0.000564222,0.0059924116,0.004243783,0.001786431,0.038417254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001733474,0.0005415753,0.009181865,0.0010375086,0.0003442613,0.0007272245,0.0013188556,0.0068213977,0.027528604,0.004620327,0.6007384,0.34540647],"study_design_scores_gemma":[0.00041525302,0.00090665615,0.0116088595,0.00028620628,0.00020821056,0.001460307,0.0015893959,0.36347038,0.05945954,0.02552066,0.5344507,0.0006238304],"about_ca_topic_score_codex":0.007227014,"about_ca_topic_score_gemma":0.011887366,"teacher_disagreement_score":0.031788997,"about_ca_system_score_codex":0.00093121437,"about_ca_system_score_gemma":0.0010026309,"threshold_uncertainty_score":0.10634482},"labels":[],"label_agreement":null},{"id":"W4400524689","doi":"10.1145/3626772.3657894","title":"LADy 💃: A Benchmark Toolkit for Latent Aspect Detection Enriched with Backtranslation Augmentation","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"Universitas Brawijaya; Natural Sciences and Engineering Research Council of Canada; Mitacs","keywords":"Benchmark (surveying); Computer science; Artificial intelligence; Cartography; Geography","score_opus":0.020828892098423927,"score_gpt":0.25270166720922943,"score_spread":0.2318727751108055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400524689","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022448493,0.004893254,0.3206362,0.0011808126,0.0008828798,0.001044547,0.10899213,0.5238666,0.016055115],"genre_scores_gemma":[0.09858629,0.0021393357,0.4578861,0.0010527688,0.00016259683,0.0022667188,0.3904091,0.035953704,0.011543354],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957474,0.0012890222,0.0005200338,0.0010432051,0.0010914867,0.0003088783],"domain_scores_gemma":[0.99251235,0.0034577749,0.00044748295,0.0020151776,0.0013179659,0.0002492849],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004933793,0.0034337007,0.0011035831,0.0054658493,0.0011366972,0.0035071196,0.003771527,0.0016580684,0.01246261],"category_scores_gemma":[0.023718333,0.0015596531,0.0028642605,0.0044678263,0.0007830081,0.005003654,0.0034383931,0.0022802413,0.017103793],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067387114,0.00027824106,0.008041078,0.0048735803,0.0007366186,0.0005153207,0.00087137154,0.025050392,0.010751254,0.0087731015,0.62520754,0.31422764],"study_design_scores_gemma":[0.00044566116,0.00034768556,0.009899788,0.00063364423,0.00028809937,0.0014115148,0.00067964743,0.4148302,0.03736584,0.026807426,0.5069341,0.00035637463],"about_ca_topic_score_codex":0.015405982,"about_ca_topic_score_gemma":0.03027849,"teacher_disagreement_score":0.015405982,"about_ca_system_score_codex":0.0016184596,"about_ca_system_score_gemma":0.003218904,"threshold_uncertainty_score":0.0416916},"labels":[],"label_agreement":null},{"id":"W4400525527","doi":"10.1145/3626772.3657675","title":"Towards Robust QA Evaluation via Open LLMs","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Universitas Brawijaya","keywords":"Computer science","score_opus":0.12426443513175954,"score_gpt":0.35075261361830107,"score_spread":0.2264881784865415,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400525527","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042499162,0.0015604335,0.8488861,0.0018342765,0.00037458993,0.0009742986,0.0037650699,0.0890638,0.011042267],"genre_scores_gemma":[0.43306774,0.00042424066,0.53838694,0.0013059054,0.0002216669,0.0015543872,0.0143175805,0.0067850384,0.00393644],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96327066,0.024292963,0.0021656128,0.0030404034,0.0064135194,0.000816857],"domain_scores_gemma":[0.9322445,0.038262688,0.0024173257,0.01340741,0.012301872,0.0013662071],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.027378952,0.0020531993,0.0013627582,0.004013799,0.0012729998,0.006869153,0.0036319294,0.0033983819,0.012069019],"category_scores_gemma":[0.13089214,0.0012156264,0.0018862652,0.0022857557,0.0020372635,0.009441689,0.0101651475,0.0046552294,0.0071735526],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021740487,0.0012404616,0.00995551,0.0023902154,0.00049419905,0.00042000442,0.0033962354,0.102209754,0.030453,0.06678805,0.11709103,0.6633875],"study_design_scores_gemma":[0.0002793253,0.00039524917,0.002095963,0.00030377752,0.0001004624,0.00020964358,0.0005693724,0.8459623,0.018712426,0.09152348,0.039711293,0.00013669414],"about_ca_topic_score_codex":0.005409557,"about_ca_topic_score_gemma":0.0059504122,"teacher_disagreement_score":0.972621,"about_ca_system_score_codex":0.0031026674,"about_ca_system_score_gemma":0.003717867,"threshold_uncertainty_score":0.14479542},"labels":[],"label_agreement":null},{"id":"W4400526199","doi":"10.1145/3626772.3657951","title":"Fine-Tuning LLaMA for Multi-Stage Text Retrieval","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":82,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Stage (stratigraphy); Information retrieval; Artificial intelligence; Natural language processing; Biology","score_opus":0.111294709817791,"score_gpt":0.3375868433020279,"score_spread":0.22629213348423688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400526199","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10187793,0.006268847,0.759969,0.00092451426,0.0005705101,0.0005846379,0.0021612335,0.121072575,0.0065707127],"genre_scores_gemma":[0.47335023,0.00085038156,0.5012764,0.001090983,0.00031991611,0.00086249196,0.007746084,0.0040411428,0.010462356],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99879646,0.00039503083,0.00009379047,0.00041033523,0.00016818423,0.00013614015],"domain_scores_gemma":[0.9978849,0.0011327062,0.000097823766,0.00044967703,0.00032566962,0.000109163295],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025456147,0.0020073461,0.0016867892,0.0015488452,0.0007544211,0.0017087392,0.0024829993,0.002049289,0.006245105],"category_scores_gemma":[0.008772771,0.00073076796,0.0018506944,0.0012025062,0.00062106096,0.0037182062,0.0016143272,0.0029241324,0.00803838],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010664159,0.0008735633,0.0043095914,0.0008984766,0.0005011359,0.00030358264,0.00033130168,0.19835913,0.056570876,0.0041820733,0.039410193,0.6931937],"study_design_scores_gemma":[0.00015077324,0.0003139586,0.00084155647,0.000032669966,0.00009915202,0.00015004614,0.00008919338,0.9620535,0.022013688,0.0041477215,0.010045464,0.000062288695],"about_ca_topic_score_codex":0.009706987,"about_ca_topic_score_gemma":0.019456958,"teacher_disagreement_score":0.009706987,"about_ca_system_score_codex":0.0011130647,"about_ca_system_score_gemma":0.002004263,"threshold_uncertainty_score":0.020891964},"labels":[],"label_agreement":null},{"id":"W4400526284","doi":"10.1145/3626772.3657942","title":"Synthetic Test Collections for Retrieval Evaluation","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"Engineering and Physical Sciences Research Council; Universitas Brawijaya","keywords":"Computer science; Test (biology); Information retrieval; Geology","score_opus":0.0516350703716649,"score_gpt":0.31526330738413105,"score_spread":0.2636282370124662,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400526284","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.59374666,0.003330553,0.33101013,0.0010249041,0.00094935915,0.009202268,0.031531185,0.009186827,0.020018132],"genre_scores_gemma":[0.6504381,0.0008750692,0.26204062,0.00077134644,0.00025782594,0.009120662,0.07009078,0.0012123514,0.0051933406],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.987329,0.0071290093,0.0011409569,0.0012019315,0.002770156,0.00042892701],"domain_scores_gemma":[0.9404209,0.02860473,0.0026791384,0.014134119,0.012847925,0.0013132562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011561751,0.0017510233,0.0012578652,0.0030844177,0.0013119759,0.0020066719,0.002884077,0.0017704472,0.0042563267],"category_scores_gemma":[0.050725985,0.00068748905,0.0013896741,0.0036693192,0.0014434473,0.0027064546,0.0028419218,0.002003505,0.0026183198],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004029228,0.0077693993,0.0384036,0.004282747,0.0010994993,0.0014024165,0.002952463,0.32473662,0.06541169,0.023912922,0.10766961,0.41832986],"study_design_scores_gemma":[0.0011908055,0.0071098283,0.027299069,0.0004933363,0.00051670196,0.001734743,0.0026003409,0.693628,0.124447465,0.021396697,0.11904223,0.00054078945],"about_ca_topic_score_codex":0.004177076,"about_ca_topic_score_gemma":0.0066844686,"teacher_disagreement_score":0.011561751,"about_ca_system_score_codex":0.0019120714,"about_ca_system_score_gemma":0.0017081137,"threshold_uncertainty_score":0.061145127},"labels":[],"label_agreement":null},{"id":"W4400526764","doi":"10.1145/3626772.3657824","title":"MTMS: Multi-teacher Multi-stage Knowledge Distillation for Reasoning-Based Machine Reading Comprehension","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Universitas Brawijaya","keywords":"Computer science; Distillation; Comprehension; Artificial intelligence; Reading (process); Reading comprehension; Natural language processing; Stage (stratigraphy); Programming language; Linguistics; Chemistry","score_opus":0.08466576282071729,"score_gpt":0.34854214489359586,"score_spread":0.26387638207287856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400526764","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018538153,0.0011485049,0.95835483,0.00067480776,0.00024055011,0.00015584007,0.0005306878,0.017349713,0.00300702],"genre_scores_gemma":[0.4384035,0.00054190547,0.5485939,0.00068666134,0.00019497865,0.00044604208,0.0023022366,0.0010819552,0.007748766],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99937457,0.00019107648,0.00003854519,0.00019792766,0.0001329551,0.00006487115],"domain_scores_gemma":[0.99897873,0.0005638722,0.000067407,0.00018596854,0.000119530494,0.000084494204],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00091070356,0.0014700559,0.0010038228,0.00057012844,0.0006214743,0.0010276992,0.002790299,0.0013986969,0.006943345],"category_scores_gemma":[0.0044240053,0.0005902481,0.0009933094,0.0005878872,0.0008744356,0.0029874016,0.0031703732,0.0031137913,0.0022917849],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064090185,0.00042683457,0.0015035743,0.0006622813,0.00016710715,0.00032312766,0.00050597783,0.23852651,0.023109104,0.022222264,0.019107487,0.6928048],"study_design_scores_gemma":[0.000057089008,0.00009256798,0.00014722922,0.000019296105,0.000018875826,0.000042896394,0.000033155193,0.9746882,0.0072864504,0.013411934,0.0041818526,0.00002049012],"about_ca_topic_score_codex":0.0051782876,"about_ca_topic_score_gemma":0.011325404,"teacher_disagreement_score":0.006943345,"about_ca_system_score_codex":0.0007977429,"about_ca_system_score_gemma":0.0020509826,"threshold_uncertainty_score":0.02322781},"labels":[],"label_agreement":null},{"id":"W4400527039","doi":"10.1109/icmi60790.2024.10585920","title":"A Deep Learning Approach for Semantic Similarity Prediction Between Question Pairs Using Siamese Network and Word Embedding Techniques","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Word embedding; Semantic similarity; Artificial intelligence; Similarity (geometry); Word (group theory); Natural language processing; Deep learning; Embedding; Linguistics","score_opus":0.03477285265856845,"score_gpt":0.2962740149750211,"score_spread":0.26150116231645265,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400527039","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11923834,0.0011997226,0.87168604,0.0007371608,0.00019699255,0.00017434025,0.00086228864,0.0030098958,0.0028952318],"genre_scores_gemma":[0.77835745,0.0006701856,0.20471749,0.0004965959,0.00020025806,0.00025698912,0.0042189704,0.000121725316,0.010960432],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996038,0.00008381235,0.00003099109,0.00016781289,0.00006719174,0.00004643072],"domain_scores_gemma":[0.99940586,0.00023285592,0.000059402697,0.000063288324,0.00020127166,0.000037268142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00067220605,0.0010264056,0.0006603905,0.0012913647,0.00034094448,0.0007173954,0.0012095164,0.0011287518,0.002366115],"category_scores_gemma":[0.001706954,0.0002896349,0.00088432536,0.0011277717,0.00042994437,0.0022803003,0.0009508079,0.0016717567,0.001045582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003687244,0.0007055184,0.006137252,0.00022947093,0.00024214988,0.00033101864,0.0004103973,0.1720078,0.021897063,0.009353057,0.012028441,0.77628917],"study_design_scores_gemma":[0.000007082886,0.00006174248,0.00044461034,0.0000067643987,0.000020152995,0.00004098961,0.000034051263,0.9917762,0.0022485552,0.0045431247,0.0008092027,0.0000075254407],"about_ca_topic_score_codex":0.006307481,"about_ca_topic_score_gemma":0.009677106,"teacher_disagreement_score":0.006307481,"about_ca_system_score_codex":0.00078240543,"about_ca_system_score_gemma":0.00081521604,"threshold_uncertainty_score":0.0125415325},"labels":[],"label_agreement":null},{"id":"W4400528755","doi":"10.1145/3626772.3657861","title":"Systematic Evaluation of Neural Retrieval Models on the Touché 2020 Argument Retrieval Subset of BEIR","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Deutsche Forschungsgemeinschaft; European Commission; HORIZON EUROPE Framework Programme; Compute Canada","keywords":"Argument (complex analysis); Computer science; Artificial intelligence; Information retrieval; Natural language processing; Biology","score_opus":0.06782064987352145,"score_gpt":0.29225979395796997,"score_spread":0.22443914408444854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400528755","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84763974,0.025481358,0.073852055,0.002184821,0.0011613432,0.0016013588,0.009755164,0.016591478,0.021732664],"genre_scores_gemma":[0.86773986,0.002138684,0.086497806,0.0010529557,0.00032356055,0.00078654767,0.03348667,0.0010017753,0.0069721453],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99082595,0.0048625483,0.0007847175,0.0015765344,0.0016358928,0.00031434908],"domain_scores_gemma":[0.9757356,0.016866095,0.00077077013,0.0035615452,0.0024692535,0.0005967069],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01590022,0.002663607,0.0022130348,0.002722945,0.00070095924,0.0022131505,0.0026963726,0.0028631794,0.003737758],"category_scores_gemma":[0.04088029,0.00064515375,0.0015577199,0.0015535769,0.0011616406,0.00408636,0.0022603031,0.0031804224,0.0030333572],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.010144269,0.003710014,0.010825074,0.0075177276,0.0027600422,0.0006513653,0.0010793916,0.20858113,0.04013593,0.003175547,0.061187137,0.65023243],"study_design_scores_gemma":[0.0010827181,0.0048700366,0.018193247,0.0004847709,0.0007048422,0.0005378682,0.0007985587,0.9170085,0.03863654,0.003998292,0.013435707,0.00024885187],"about_ca_topic_score_codex":0.0060278396,"about_ca_topic_score_gemma":0.0071744802,"teacher_disagreement_score":0.01590022,"about_ca_system_score_codex":0.0017805955,"about_ca_system_score_gemma":0.0011550803,"threshold_uncertainty_score":0.08408946},"labels":[],"label_agreement":null},{"id":"W4400528870","doi":"10.1145/3626772.3657992","title":"LLM4Eval: Large Language Model for Evaluation in IR","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); University of Waterloo","funders":"Engineering and Physical Sciences Research Council; Vrije Universiteit Amsterdam; Universiteit van Amsterdam","keywords":"Computer science; Programming language; Natural language processing","score_opus":0.056575431006054715,"score_gpt":0.3574910327463317,"score_spread":0.300915601740277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400528870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017193507,0.0036614998,0.8489829,0.0035124396,0.0022556812,0.0030515702,0.010714946,0.0974691,0.01315841],"genre_scores_gemma":[0.19518794,0.0007569673,0.74957216,0.002394988,0.00043431637,0.005549413,0.027762549,0.011281926,0.007059746],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.93428284,0.053303156,0.0023049412,0.0033640144,0.0059215883,0.00082349376],"domain_scores_gemma":[0.9314945,0.04927431,0.0012636159,0.0113543,0.004877843,0.0017353656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.053489514,0.0030936145,0.0024729783,0.002995404,0.001638736,0.006747421,0.00672764,0.0047547757,0.02127181],"category_scores_gemma":[0.111072965,0.0015594696,0.0033414855,0.0018433366,0.0015995296,0.007882156,0.008430813,0.007467748,0.00828849],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039788675,0.0013402574,0.0044136513,0.0027298492,0.0018337717,0.00047218715,0.0010786296,0.081303656,0.006516882,0.028332338,0.32767475,0.5403252],"study_design_scores_gemma":[0.0012918093,0.001016485,0.0018179015,0.00042751155,0.00027818987,0.00027715284,0.00035604037,0.85513127,0.009861694,0.06343696,0.06586042,0.00024445876],"about_ca_topic_score_codex":0.006608891,"about_ca_topic_score_gemma":0.012150301,"teacher_disagreement_score":0.053489514,"about_ca_system_score_codex":0.0032570176,"about_ca_system_score_gemma":0.0038681796,"threshold_uncertainty_score":0.282883},"labels":[],"label_agreement":null},{"id":"W4400650325","doi":"10.1145/3678003","title":"A Knowledge Graph Embedding Model for Answering Factoid Entity Questions","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Question answering; Computer science; Embedding; Knowledge graph; Information retrieval; Graph; Artificial intelligence; Natural language processing; Theoretical computer science","score_opus":0.03891582147792144,"score_gpt":0.30215229729600357,"score_spread":0.2632364758180821,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400650325","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014839797,0.0014391423,0.9747021,0.0010579387,0.00011245337,0.0002073453,0.0027306203,0.0019725314,0.0029380375],"genre_scores_gemma":[0.39432955,0.0022377463,0.5692882,0.0009820805,0.00028774457,0.00059368886,0.021135224,0.00029615915,0.01084961],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925274,0.00021478639,0.00005494115,0.00027197026,0.00015952272,0.00004608362],"domain_scores_gemma":[0.9987155,0.00077125814,0.000089481924,0.00016500159,0.00021571263,0.000043000604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00082469743,0.00088588125,0.00058817276,0.0018240544,0.00041691953,0.0010885074,0.001304488,0.0016139515,0.0033395914],"category_scores_gemma":[0.0047013992,0.00032095914,0.001238884,0.0018394592,0.00044439247,0.004218193,0.0010879547,0.0016797787,0.0015703809],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036849637,0.00060174544,0.004557539,0.0009338883,0.000302885,0.00048620562,0.0009628263,0.2510891,0.0095305685,0.077862814,0.043365028,0.6099389],"study_design_scores_gemma":[0.000021256923,0.00006433253,0.0006543196,0.000042709944,0.000053961037,0.00015789537,0.00013000598,0.9210511,0.0014448133,0.06502626,0.011329757,0.000023554594],"about_ca_topic_score_codex":0.0080407,"about_ca_topic_score_gemma":0.014171173,"teacher_disagreement_score":0.0080407,"about_ca_system_score_codex":0.00089460285,"about_ca_system_score_gemma":0.00094291335,"threshold_uncertainty_score":0.015987813},"labels":[],"label_agreement":null},{"id":"W4400680797","doi":"10.1109/saner60148.2024.00017","title":"Gloss: Guiding Large Language Models to Answer Questions from System Logs","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"National Natural Science Foundation of China","keywords":"Gloss (optics); Computer science; Natural language processing; Programming language; Artificial intelligence","score_opus":0.03842013210703462,"score_gpt":0.27622549928726964,"score_spread":0.237805367180235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400680797","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04564347,0.001279587,0.87934136,0.0019120544,0.00029430081,0.0011345013,0.008129243,0.059253942,0.0030115282],"genre_scores_gemma":[0.30460954,0.00049198145,0.6503419,0.002249154,0.00022738114,0.0018603151,0.034100037,0.0020716907,0.004048051],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99665904,0.001636417,0.0002488516,0.0009267013,0.00037097489,0.00015800857],"domain_scores_gemma":[0.98641485,0.010762776,0.0003526416,0.0011461914,0.0010525092,0.00027102785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004956426,0.0030703638,0.0011912794,0.0019796076,0.00088396855,0.0023725634,0.003411683,0.0029057483,0.0055252733],"category_scores_gemma":[0.02296462,0.00093219936,0.002311577,0.0009288277,0.001052536,0.0048842514,0.0030186882,0.0042115664,0.0038117445],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013951106,0.0014392342,0.013661939,0.0020955938,0.00043742554,0.0008473847,0.0036492965,0.16752627,0.025246108,0.013474216,0.083523974,0.6867035],"study_design_scores_gemma":[0.00013773468,0.0001451607,0.00069864775,0.000053451673,0.00005585124,0.00012522802,0.00038552246,0.97011393,0.0049247276,0.014311683,0.009004887,0.00004306718],"about_ca_topic_score_codex":0.011454955,"about_ca_topic_score_gemma":0.020801624,"teacher_disagreement_score":0.011454955,"about_ca_system_score_codex":0.001673644,"about_ca_system_score_gemma":0.0022481338,"threshold_uncertainty_score":0.026212394},"labels":[],"label_agreement":null},{"id":"W4400909682","doi":"10.1109/icde60146.2024.00466","title":"Towards Explainability in Retrieval-Augmented LLMs","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Waterloo","funders":"","keywords":"Computer science","score_opus":0.025830372725308248,"score_gpt":0.2820160440700754,"score_spread":0.25618567134476716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400909682","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01615463,0.00022985463,0.9776354,0.002370462,0.000035147335,0.00011081365,0.00026538558,0.0014933273,0.0017050505],"genre_scores_gemma":[0.38263214,0.00038665836,0.6115509,0.000739828,0.00017872294,0.00033771966,0.0010450116,0.00071694923,0.0024120123],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9896539,0.0059903124,0.0007339113,0.0016082183,0.0015926616,0.00042099386],"domain_scores_gemma":[0.9053929,0.07087722,0.005189118,0.014078492,0.0037929872,0.0006692745],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014508986,0.0010143099,0.0009533827,0.0038568964,0.0015686144,0.0054869945,0.0022773622,0.0027225385,0.005708608],"category_scores_gemma":[0.09817782,0.00090894726,0.0029144983,0.0020418882,0.0046824226,0.014376943,0.008504863,0.004028449,0.0009338177],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016057953,0.00008838425,0.0061593154,0.0004928573,0.00014900122,0.0006978407,0.009939735,0.040963147,0.0042269463,0.8371183,0.0029746352,0.09702912],"study_design_scores_gemma":[0.000036606878,0.000045629607,0.0005792989,0.00015123129,0.000083691666,0.00022321843,0.0006609488,0.17341636,0.004380852,0.804262,0.016112544,0.000047646667],"about_ca_topic_score_codex":0.0033373733,"about_ca_topic_score_gemma":0.00311079,"teacher_disagreement_score":0.014508986,"about_ca_system_score_codex":0.002110462,"about_ca_system_score_gemma":0.0021908372,"threshold_uncertainty_score":0.0767318},"labels":[],"label_agreement":null},{"id":"W4400910582","doi":"10.1109/ic3se62002.2024.10593288","title":"Development of Language Model on Biomedical Domain to Pretrain Natural Language Processing","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Horizon College and Seminary","funders":"","keywords":"Computer science; Natural language processing; Domain (mathematical analysis); Natural language; Human–computer interaction; Artificial intelligence","score_opus":0.015137226554105562,"score_gpt":0.29071069471860417,"score_spread":0.2755734681644986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400910582","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040896114,0.0009562238,0.91951317,0.0018436126,0.0004678121,0.00037848775,0.004413147,0.023002787,0.008528617],"genre_scores_gemma":[0.42447698,0.0010798632,0.5158377,0.0015480354,0.00023641025,0.0009575255,0.030730944,0.0015711962,0.023561383],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99951994,0.0001252602,0.00003601539,0.00016275753,0.00010231485,0.000053703992],"domain_scores_gemma":[0.99889743,0.00045522174,0.000044725668,0.0001529229,0.0003876297,0.00006203488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001024182,0.0008609118,0.00054604164,0.00064641057,0.00045812593,0.0009862521,0.0014512665,0.0009741167,0.006394559],"category_scores_gemma":[0.0029673497,0.0004954283,0.0011930787,0.0005499578,0.0002642784,0.0022299085,0.0008247187,0.0027672853,0.0049547306],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003753287,0.00047158994,0.0046872054,0.0004857598,0.00023594333,0.00046342882,0.0003420703,0.35989657,0.034515563,0.017390182,0.068430685,0.5127057],"study_design_scores_gemma":[0.000019242261,0.000064758766,0.0005364126,0.000029145029,0.000031031323,0.00009015067,0.00005522143,0.9750004,0.007939951,0.0067378934,0.009476071,0.00001963949],"about_ca_topic_score_codex":0.012200596,"about_ca_topic_score_gemma":0.013937238,"teacher_disagreement_score":0.012200596,"about_ca_system_score_codex":0.00095157634,"about_ca_system_score_gemma":0.0018345785,"threshold_uncertainty_score":0.02425915},"labels":[],"label_agreement":null},{"id":"W4400949264","doi":"10.1038/s41586-024-07566-y","title":"AI models collapse when trained on recursively generated data","year":2024,"lang":"en","type":"article","venue":"Nature","topic":"Topic Modeling","field":"Computer Science","cited_by":583,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Generative grammar; Generative model; Intuition; Computer science; Variety (cybernetics); The Internet; Artificial intelligence; Data science; Psychology; Cognitive science; World Wide Web","score_opus":0.0509041440660831,"score_gpt":0.30053330020218677,"score_spread":0.24962915613610367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400949264","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29190364,0.0030118134,0.66898596,0.0028792282,0.00074053905,0.00023571173,0.001571375,0.010917766,0.019753927],"genre_scores_gemma":[0.86622876,0.0005348307,0.11904196,0.00074392307,0.00011592137,0.0001681255,0.0025619618,0.0005013198,0.010103171],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99920136,0.00029497608,0.000046110614,0.00028274796,0.00010868985,0.00006620877],"domain_scores_gemma":[0.99614304,0.0023685663,0.000121965495,0.00080728176,0.00045781553,0.000101399935],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0027221746,0.00084983214,0.0006725438,0.00070675946,0.00046520063,0.0012895297,0.0013201425,0.0014930619,0.0032097888],"category_scores_gemma":[0.010986844,0.00069609034,0.0011634778,0.0006535237,0.0007109308,0.0019797583,0.0012468351,0.0035300828,0.002546047],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018184203,0.00018253164,0.0075658453,0.00018678326,0.00030350868,0.00012224002,0.00034874782,0.7034908,0.0047470434,0.016215773,0.010242931,0.2564119],"study_design_scores_gemma":[0.000005852139,0.000028782008,0.0005506502,0.000012368602,0.000011693596,0.000022924094,0.000018280309,0.99165887,0.0013916777,0.0049641533,0.0013274017,0.0000074244144],"about_ca_topic_score_codex":0.01773535,"about_ca_topic_score_gemma":0.023128645,"teacher_disagreement_score":0.9972778,"about_ca_system_score_codex":0.0014912006,"about_ca_system_score_gemma":0.0012640336,"threshold_uncertainty_score":0.035264254},"labels":[],"label_agreement":null},{"id":"W4400977996","doi":"10.1016/j.patter.2024.101030","title":"Exploring the reversal curse and other deductive logical reasoning in BERT and GPT-based large language models","year":2024,"lang":"en","type":"article","venue":"Patterns","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; CHEO Research Institute","keywords":"Transformer; Computer science; Encoder; Curse; Generative grammar; Logical reasoning; Comprehension; Artificial intelligence; Natural language processing; Autoencoder; Intersection (aeronautics); Context (archaeology); Machine learning; Programming language; Deep learning","score_opus":0.07904218656723885,"score_gpt":0.2861965466404941,"score_spread":0.20715436007325524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400977996","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15733956,0.00020491135,0.834919,0.0011922735,0.00003255991,0.000081866376,0.00023158044,0.0012001265,0.004798195],"genre_scores_gemma":[0.8239076,0.00015143373,0.17335236,0.00022960603,0.000018490095,0.00008693537,0.0004432674,0.00017062122,0.0016396416],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998544,0.0008052455,0.00007702267,0.0003015272,0.00020009145,0.00007221285],"domain_scores_gemma":[0.9838193,0.01342876,0.00057819556,0.0013838349,0.0005398439,0.00025009774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004582939,0.0007280914,0.00050168566,0.00067448773,0.00043898425,0.002550506,0.0014313994,0.0010378136,0.0026328405],"category_scores_gemma":[0.028841268,0.00067581545,0.0008151943,0.00064774725,0.0014498159,0.006960858,0.0018717893,0.0033302498,0.0004641646],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038624625,0.00028360393,0.012912973,0.00045628354,0.00020343637,0.0008131923,0.004552391,0.37153095,0.013898273,0.3117998,0.0034262694,0.27973652],"study_design_scores_gemma":[0.00001679497,0.000042257518,0.00045159808,0.000020375106,0.000021121698,0.00010189127,0.00017125806,0.8600039,0.0030588147,0.13491832,0.0011735944,0.000020069188],"about_ca_topic_score_codex":0.0051800124,"about_ca_topic_score_gemma":0.010408068,"teacher_disagreement_score":0.0051800124,"about_ca_system_score_codex":0.0013600761,"about_ca_system_score_gemma":0.0014854423,"threshold_uncertainty_score":0.024237216},"labels":[],"label_agreement":null},{"id":"W4401023522","doi":"10.24963/ijcai.2024/714","title":"MASTER: A Multi-granularity Invariant Structure Clustering Scheme for Multi-view Clustering","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Samsung; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Natural language processing; Memorization; Artificial intelligence; Programming language; Linguistics","score_opus":0.08915198147682553,"score_gpt":0.3062386796759296,"score_spread":0.2170866981991041,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401023522","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006885484,0.0003915441,0.9894756,0.00010255824,0.000046214984,0.00007701047,0.00022938821,0.0020891824,0.0007029157],"genre_scores_gemma":[0.16683555,0.00042756132,0.8243868,0.00027444196,0.00011304787,0.00021607285,0.0026923637,0.00057787483,0.004476324],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985801,0.0002448455,0.00007902521,0.00045549573,0.00047424514,0.00016625873],"domain_scores_gemma":[0.9985397,0.00017596458,0.00015262196,0.00063105905,0.00037973167,0.000120920566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016803058,0.0013078245,0.0020545286,0.0023813883,0.0011932958,0.001566439,0.0038113245,0.001781227,0.0023448628],"category_scores_gemma":[0.003640403,0.0006562443,0.0018511668,0.0025542355,0.00092924543,0.0029726303,0.0032302176,0.0022770977,0.002014857],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037677633,0.00024030604,0.002871354,0.00025054903,0.0003508127,0.00011686779,0.0005040853,0.20236358,0.027632192,0.033855215,0.022909362,0.7085289],"study_design_scores_gemma":[0.000027032966,0.00006238266,0.0005626027,0.000011679243,0.000025808986,0.00008295321,0.00006172824,0.9747971,0.005221844,0.01403946,0.005072279,0.00003516891],"about_ca_topic_score_codex":0.00735255,"about_ca_topic_score_gemma":0.013155152,"teacher_disagreement_score":0.00735255,"about_ca_system_score_codex":0.0016423693,"about_ca_system_score_gemma":0.0016555776,"threshold_uncertainty_score":0.014619529},"labels":[],"label_agreement":null},{"id":"W4401023556","doi":"10.24963/ijcai.2024/890","title":"DiffECG: Diffusion Model-Powered Label-Efficient and Personalized Arrhythmia Diagnosis","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":212,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Division of Chemistry; National Science Foundation","keywords":"Computer science; Programming language; Data science","score_opus":0.022838901854467176,"score_gpt":0.25826486838301205,"score_spread":0.23542596652854486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401023556","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009441853,0.00040333325,0.9825979,0.00026108566,0.000091268696,0.0000869959,0.00023435193,0.0060964,0.0007867539],"genre_scores_gemma":[0.32207763,0.00056226953,0.66527313,0.00072611345,0.0002475148,0.00027219125,0.0018925053,0.0012206697,0.0077279895],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992229,0.00014870806,0.000033766002,0.00031354817,0.00021276556,0.000068440975],"domain_scores_gemma":[0.9991184,0.0002694467,0.000094545634,0.00029308247,0.00015074268,0.000073691226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013579263,0.0011721713,0.0012840617,0.0012374441,0.0004999971,0.0008174501,0.00213457,0.0017038041,0.0016915399],"category_scores_gemma":[0.0033055127,0.0004495622,0.0011402707,0.00089318474,0.00064041035,0.0013584296,0.0019996825,0.0018121877,0.001161899],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032277184,0.0002746355,0.0031477374,0.00017584617,0.00021871935,0.00025200238,0.00019449089,0.13509876,0.028688282,0.0072852857,0.023582125,0.8007594],"study_design_scores_gemma":[0.00003428533,0.00004967133,0.0005309826,0.000010327815,0.000023176988,0.00018641117,0.000018317081,0.98111224,0.0074670254,0.007260074,0.0032816848,0.000025844161],"about_ca_topic_score_codex":0.005015024,"about_ca_topic_score_gemma":0.00881467,"teacher_disagreement_score":0.005015024,"about_ca_system_score_codex":0.000839556,"about_ca_system_score_gemma":0.0011524282,"threshold_uncertainty_score":0.009971678},"labels":[],"label_agreement":null},{"id":"W4401023643","doi":"10.24963/ijcai.2024/917","title":"DFMU: Distribution-based Framework for Modeling Aleatoric Uncertainty in Multimodal Sentiment Analysis","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); Vector Institute; Western University","funders":"","keywords":"Computer science; Context (archaeology); Length measurement; Geology; Physics","score_opus":0.022695320868185778,"score_gpt":0.2945012643199514,"score_spread":0.2718059434517656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401023643","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007444656,0.00043479714,0.9903634,0.0003185841,0.00003589205,0.000049005124,0.00022321714,0.00036360344,0.00076675665],"genre_scores_gemma":[0.6448602,0.00096707104,0.34770197,0.00070137624,0.00036356237,0.00047081505,0.0014345846,0.00032124898,0.0031791958],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9985783,0.0005996327,0.00008338007,0.00036929763,0.00025336278,0.000116135736],"domain_scores_gemma":[0.99717927,0.0019787003,0.00025276473,0.00015804266,0.00034511875,0.0000862251],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038347016,0.0012972765,0.0011820394,0.001608847,0.00057366665,0.0015551902,0.001898928,0.0011614187,0.0016746775],"category_scores_gemma":[0.009816165,0.00057085656,0.0011416739,0.0011720334,0.0010670867,0.002492078,0.0017567974,0.0022872135,0.0004743661],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021377511,0.000095235235,0.0046678,0.0002281892,0.00023384707,0.00016257571,0.00041081788,0.7611115,0.004096325,0.046594456,0.0058527817,0.17633267],"study_design_scores_gemma":[0.0000046177342,0.000016015734,0.00024117483,0.000010757164,0.0000102685235,0.000015393107,0.000014797196,0.9840853,0.000288679,0.01458203,0.0007222167,0.000008764959],"about_ca_topic_score_codex":0.007826968,"about_ca_topic_score_gemma":0.010287025,"teacher_disagreement_score":0.007826968,"about_ca_system_score_codex":0.0017556662,"about_ca_system_score_gemma":0.0013291821,"threshold_uncertainty_score":0.020280063},"labels":[],"label_agreement":null},{"id":"W4401024755","doi":"10.24963/ijcai.2024/634","title":"A Primal-dual Perspective for Distributed TD-learning","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"National Natural Science Foundation of China; Tencent","keywords":"Inference; Computer science; Natural language processing; Language model; Artificial intelligence","score_opus":0.020537589230253894,"score_gpt":0.2842636548327588,"score_spread":0.2637260656025049,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401024755","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029209869,0.00023830171,0.9949084,0.00028922196,0.0000441125,0.000014482339,0.000016788596,0.000019348894,0.0015482871],"genre_scores_gemma":[0.6984452,0.0010753262,0.29169038,0.00041273213,0.0002839167,0.0002825327,0.0001118582,0.00009738876,0.0076007126],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99941075,0.00027608612,0.000023639019,0.00011230525,0.00012608636,0.00005118856],"domain_scores_gemma":[0.99781066,0.0015258793,0.00016493465,0.00010131004,0.00029725378,0.00009994018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026264288,0.00094999745,0.0011272765,0.00041313292,0.00036540796,0.0012806777,0.0013312573,0.0012306671,0.002778665],"category_scores_gemma":[0.005701737,0.00041374838,0.0006933272,0.00057974434,0.0014632313,0.0014584576,0.0015043383,0.0023262403,0.00026449267],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046405956,0.000041525927,0.0003160538,0.0001034498,0.000032205895,0.00005001346,0.000049993385,0.84126246,0.0006046244,0.14569786,0.0007250389,0.011070324],"study_design_scores_gemma":[0.0000041014196,0.000011702927,0.000014382441,0.0000043791906,0.0000025942531,0.000007011006,0.000004613743,0.9815933,0.000100791534,0.017971102,0.00028372608,0.0000022476218],"about_ca_topic_score_codex":0.0019972734,"about_ca_topic_score_gemma":0.0012647375,"teacher_disagreement_score":0.002778665,"about_ca_system_score_codex":0.0011383161,"about_ca_system_score_gemma":0.0013270943,"threshold_uncertainty_score":0.013890088},"labels":[],"label_agreement":null},{"id":"W4401042062","doi":"10.18653/v1/2024.semeval-1.265","title":"Edinburgh Clinical NLP at SemEval-2024 Task 2: Fine-tune your model unless you have access to GPT-4","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Engineering and Physical Sciences Research Council; National Institute for Health and Care Research; Canadian Institute of Steel Construction; UK Research and Innovation; Nvidia; Accenture; Cisco Systems","keywords":"SemEval; Task (project management); Computer science; Natural language processing; Artificial intelligence; Speech recognition; Engineering","score_opus":0.13986571912916657,"score_gpt":0.4029602116996834,"score_spread":0.2630944925705168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042062","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24356209,0.017138246,0.3037397,0.032640625,0.003807662,0.0038619314,0.25361377,0.10789519,0.03374076],"genre_scores_gemma":[0.56894493,0.0015532651,0.19865413,0.005535336,0.0007054276,0.00233454,0.2048334,0.0053766896,0.01206227],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98813397,0.0076371995,0.000770614,0.002136829,0.0010590624,0.00026232487],"domain_scores_gemma":[0.9510677,0.040199377,0.0011622291,0.0047505214,0.0018270224,0.0009931951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027071238,0.0023592913,0.0019588983,0.0017516003,0.00081747817,0.0028598814,0.0031605044,0.004077689,0.019867294],"category_scores_gemma":[0.10048825,0.0010432756,0.0030194318,0.001524623,0.00082963624,0.0032090947,0.0026475503,0.00397353,0.010300033],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00878991,0.00096821244,0.027611135,0.005901439,0.0028123907,0.001378387,0.0005828979,0.13079393,0.008511049,0.0056530572,0.43655023,0.3704473],"study_design_scores_gemma":[0.0076627475,0.0025372712,0.019798733,0.0009311268,0.0015896641,0.002530341,0.00044201635,0.72398597,0.022353588,0.032244414,0.1854538,0.00047042573],"about_ca_topic_score_codex":0.011458611,"about_ca_topic_score_gemma":0.016797218,"teacher_disagreement_score":0.027071238,"about_ca_system_score_codex":0.0022298766,"about_ca_system_score_gemma":0.0045094606,"threshold_uncertainty_score":0.14316809},"labels":[],"label_agreement":null},{"id":"W4401042103","doi":"10.18653/v1/2024.semeval-1.79","title":"TLDR at SemEval-2024 Task 2: T5-generated clinical-Language summaries for DeBERTa Report Analysis","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"SemEval; Computer science; Natural language processing; Task (project management); Artificial intelligence; Information retrieval; Engineering","score_opus":0.043895534026346196,"score_gpt":0.35817205584752815,"score_spread":0.31427652182118193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042103","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054525807,0.0012889192,0.7782688,0.0038355927,0.0010071876,0.0018631037,0.036810525,0.110658154,0.01174185],"genre_scores_gemma":[0.23974445,0.00038711578,0.70842594,0.000942428,0.00031761904,0.001241644,0.03969288,0.005075141,0.0041728],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9906221,0.0052671228,0.0008553514,0.0016954091,0.0013623566,0.0001976849],"domain_scores_gemma":[0.9448137,0.043135855,0.002145672,0.0053511383,0.0037613318,0.0007922102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011047158,0.0018488457,0.0008511425,0.0017109124,0.0007078288,0.0025661087,0.002521741,0.0025632118,0.02306193],"category_scores_gemma":[0.069311,0.0006773253,0.0017203318,0.00080456876,0.000712316,0.0031691138,0.0042360807,0.0033492232,0.010619054],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040606866,0.00079680065,0.0063987407,0.005572938,0.0006498357,0.0024807677,0.004334827,0.04353935,0.05751991,0.030887326,0.24135979,0.6023991],"study_design_scores_gemma":[0.0016131988,0.0012556022,0.0043225614,0.00068027206,0.00045028335,0.0026111563,0.0016038154,0.5391364,0.13327706,0.072302386,0.24238901,0.00035821777],"about_ca_topic_score_codex":0.0019952357,"about_ca_topic_score_gemma":0.0027344197,"teacher_disagreement_score":0.02306193,"about_ca_system_score_codex":0.0012863862,"about_ca_system_score_gemma":0.0033607685,"threshold_uncertainty_score":0.07714987},"labels":[],"label_agreement":null},{"id":"W4401042217","doi":"10.18653/v1/2024.naacl-industry.33","title":"Tiny Titans: Can Smaller Large Language Models Punch Above Their Weight in the Real World for Meeting Summarization?","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Stornoway Diamond (Canada)","funders":"","keywords":"Automatic summarization; Computer science; Artificial intelligence","score_opus":0.02476747805012692,"score_gpt":0.26570334256592837,"score_spread":0.24093586451580146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042217","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062031988,0.013322097,0.8089259,0.039727606,0.006130704,0.00041283492,0.0060288794,0.041940033,0.021479951],"genre_scores_gemma":[0.51116425,0.0038814922,0.43916032,0.0058347587,0.0033329881,0.0006039257,0.011627027,0.0050048465,0.019390376],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99808455,0.00095768226,0.0001185675,0.00040470134,0.00031464922,0.000119730605],"domain_scores_gemma":[0.9910761,0.0053688483,0.0003685561,0.0017764515,0.0010023474,0.0004075819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048129857,0.0017475653,0.0017262101,0.0014971445,0.0014551712,0.0037364566,0.0028165155,0.002136833,0.016974628],"category_scores_gemma":[0.032276064,0.0010080496,0.0009789122,0.0014365236,0.001064949,0.015801735,0.0026353623,0.0031427687,0.014773583],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030734332,0.00027728075,0.0031185206,0.0006694582,0.00037691035,0.00022232451,0.0013156593,0.023675341,0.008131537,0.03413358,0.20555346,0.71945244],"study_design_scores_gemma":[0.0006791994,0.0005597133,0.0016390637,0.00019591968,0.00027499482,0.00024407492,0.001341947,0.66414034,0.0078464085,0.21474084,0.10819987,0.00013759638],"about_ca_topic_score_codex":0.00457449,"about_ca_topic_score_gemma":0.012065454,"teacher_disagreement_score":0.016974628,"about_ca_system_score_codex":0.00088210567,"about_ca_system_score_gemma":0.0012589942,"threshold_uncertainty_score":0.056785762},"labels":[],"label_agreement":null},{"id":"W4401042328","doi":"10.18653/v1/2024.naacl-long.30","title":"DuRE: Dual Contrastive Self Training for Semi-Supervised Relation Extraction","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Dual (grammatical number); Computer science; Relationship extraction; Relation (database); Artificial intelligence; Training (meteorology); Extraction (chemistry); Natural language processing; Pattern recognition (psychology); Data mining; Chromatography; Linguistics; Chemistry","score_opus":0.04646948887432286,"score_gpt":0.289196637814587,"score_spread":0.24272714894026412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042328","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022609867,0.0015338797,0.88548064,0.0003744313,0.0004547755,0.00033990873,0.0042374604,0.0798871,0.005081808],"genre_scores_gemma":[0.17369722,0.00049815356,0.77908474,0.0006272644,0.0002527204,0.00081035047,0.027020924,0.0038857171,0.014123015],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998185,0.0005130409,0.00013846983,0.0007601441,0.00028219246,0.00012103069],"domain_scores_gemma":[0.9964276,0.0020529574,0.00012219351,0.0008468817,0.00043385074,0.000116598414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026439643,0.0017981841,0.0012869745,0.0023609383,0.0010775446,0.001952049,0.0032496217,0.002173614,0.01230992],"category_scores_gemma":[0.0056563728,0.0011575508,0.0015298707,0.0015717744,0.00074191624,0.00411137,0.002849666,0.002902207,0.011884247],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087159325,0.0005167556,0.0020934092,0.0005575747,0.00032189986,0.00029721114,0.0004062617,0.013605209,0.03307139,0.0061412235,0.091659606,0.8504579],"study_design_scores_gemma":[0.00021364844,0.00037720596,0.0022731104,0.00010949633,0.00013732056,0.0004544567,0.00023812486,0.89088446,0.042431507,0.018750127,0.044045713,0.00008489341],"about_ca_topic_score_codex":0.0026387346,"about_ca_topic_score_gemma":0.00946822,"teacher_disagreement_score":0.01230992,"about_ca_system_score_codex":0.0006322972,"about_ca_system_score_gemma":0.0012168126,"threshold_uncertainty_score":0.04118073},"labels":[],"label_agreement":null},{"id":"W4401042510","doi":"10.18653/v1/2024.findings-naacl.58","title":"GraSAME: Injecting Token-Level Structural Information to Pretrained Language Models via Graph-guided Self-Attention Mechanism","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Security token; Mechanism (biology); Language model; Graph; Artificial intelligence; Theoretical computer science; Computer network","score_opus":0.023355761985992582,"score_gpt":0.2635152800684721,"score_spread":0.24015951808247948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042510","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02292349,0.0006606422,0.9323844,0.0006249767,0.00034062256,0.00016904225,0.0010466301,0.038350474,0.003499817],"genre_scores_gemma":[0.4784771,0.00068250304,0.48641753,0.002117262,0.00022900834,0.0005974766,0.0069913743,0.0036394869,0.020848159],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995504,0.00011988952,0.000019414418,0.00019242596,0.00007348169,0.00004431659],"domain_scores_gemma":[0.99886715,0.00055353367,0.000059613565,0.00026123997,0.00019701832,0.00006135422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00089533755,0.0020245856,0.0008032807,0.00081288855,0.00035074225,0.0008577323,0.0028972356,0.0013793495,0.0069508846],"category_scores_gemma":[0.003742587,0.0007232526,0.0013843583,0.0006482274,0.0006379937,0.0037672815,0.002144509,0.0027277584,0.0042799865],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042483435,0.0003893999,0.0017870509,0.00047566125,0.000325991,0.00047152053,0.0003706535,0.26244053,0.042348623,0.016053958,0.044168647,0.6307432],"study_design_scores_gemma":[0.000028085946,0.00006618152,0.00018773109,0.000014061032,0.000037062735,0.00005171958,0.000030202253,0.97854006,0.0074866726,0.009869133,0.0036703139,0.000018718449],"about_ca_topic_score_codex":0.0059710653,"about_ca_topic_score_gemma":0.0155242365,"teacher_disagreement_score":0.0069508846,"about_ca_system_score_codex":0.0007085928,"about_ca_system_score_gemma":0.0009651694,"threshold_uncertainty_score":0.023253024},"labels":[],"label_agreement":null},{"id":"W4401042714","doi":"10.18653/v1/2024.naacl-srw.5","title":"SMARTR: A Framework for Early Detection using Survival Analysis of Longitudinal Texts","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Computer science; Natural language processing","score_opus":0.057352196951382446,"score_gpt":0.3218239795214972,"score_spread":0.2644717825701148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042714","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024425562,0.0005407324,0.9843828,0.00022693102,0.0000734357,0.00008568514,0.001200504,0.010721534,0.0003257787],"genre_scores_gemma":[0.08734042,0.00096457166,0.8965522,0.00026386764,0.0004340533,0.0008132967,0.0070263073,0.0019020733,0.0047031934],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971033,0.0013093736,0.00017199617,0.00069690286,0.0005471415,0.00017131337],"domain_scores_gemma":[0.9862518,0.0094768545,0.00089163176,0.0015774991,0.0014060673,0.0003961663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010232687,0.001900391,0.0019739254,0.0075421026,0.0010070753,0.0027583293,0.0032384743,0.0022730203,0.0074565774],"category_scores_gemma":[0.023527835,0.0013879429,0.0026017446,0.0037462818,0.00085836597,0.0034676755,0.0031670867,0.002854112,0.007968909],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008122687,0.0004098644,0.015925502,0.0006282664,0.0007351247,0.00050447567,0.00087877887,0.057700533,0.0073992233,0.030684669,0.05429739,0.83002394],"study_design_scores_gemma":[0.00009996605,0.00014703968,0.0026824716,0.00009246953,0.00013015246,0.00023858111,0.00012928167,0.90410054,0.0032088882,0.07082269,0.01825959,0.00008832883],"about_ca_topic_score_codex":0.0060909074,"about_ca_topic_score_gemma":0.009381713,"teacher_disagreement_score":0.010232687,"about_ca_system_score_codex":0.000696762,"about_ca_system_score_gemma":0.001829091,"threshold_uncertainty_score":0.05411625},"labels":[],"label_agreement":null},{"id":"W4401043383","doi":"10.18653/v1/2024.clinicalnlp-1.49","title":"Edinburgh Clinical NLP at MEDIQA-CORR 2024: Guiding Large Language Models with Hints","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Engineering and Physical Sciences Research Council; National Institute for Health and Care Research; Canadian Institute of Steel Construction; UK Research and Innovation; Nvidia; Accenture; Cisco Systems","keywords":"Natural language processing; Computer science; Artificial intelligence; Linguistics; Philosophy","score_opus":0.06622473689383701,"score_gpt":0.334065904672627,"score_spread":0.26784116777878997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401043383","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0185278,0.0015346075,0.76348263,0.014930034,0.0010195739,0.0008806825,0.068484634,0.117852606,0.013287502],"genre_scores_gemma":[0.2644695,0.0008358918,0.6232252,0.0033468995,0.0005272158,0.00086586614,0.07583637,0.016083797,0.014809284],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99511814,0.0029076533,0.00026099864,0.0009555688,0.00056132494,0.00019632324],"domain_scores_gemma":[0.98093075,0.0151436,0.0003136718,0.0013181642,0.0016932457,0.0006006017],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006139224,0.0017511976,0.0010970901,0.0015209971,0.0010217665,0.00410332,0.0019871576,0.0029362708,0.05003808],"category_scores_gemma":[0.04577396,0.0012446565,0.0022512572,0.001145414,0.0006923292,0.003734205,0.002922045,0.0036370337,0.024821192],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020685468,0.0003448854,0.007907934,0.0017499307,0.00036996248,0.0006473522,0.0014246183,0.05056933,0.007344491,0.027164618,0.56767964,0.33272856],"study_design_scores_gemma":[0.000661774,0.00022286127,0.0014526638,0.00038404102,0.00031397215,0.00055216905,0.00071964174,0.71779704,0.012884518,0.10110902,0.16376683,0.00013556918],"about_ca_topic_score_codex":0.016462337,"about_ca_topic_score_gemma":0.026443707,"teacher_disagreement_score":0.05003808,"about_ca_system_score_codex":0.0018346005,"about_ca_system_score_gemma":0.0052676117,"threshold_uncertainty_score":0.16739404},"labels":[],"label_agreement":null},{"id":"W4401043387","doi":"10.18653/v1/2024.semeval-1.188","title":"BD-NLP at SemEval-2024 Task 2: Investigating Generative and Discriminative Models for Clinical Inference with Knowledge Augmentation","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Discriminative model; SemEval; Artificial intelligence; Computer science; Generative grammar; Inference; Natural language processing; Task (project management); Generative model; Engineering","score_opus":0.17368123730901736,"score_gpt":0.41053544762308986,"score_spread":0.2368542103140725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401043387","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22235239,0.014584953,0.61071014,0.020203331,0.0019065207,0.0026336669,0.046766628,0.059221312,0.021621013],"genre_scores_gemma":[0.5514298,0.0011766577,0.37124923,0.0041099167,0.0005891741,0.0012560755,0.061103374,0.0012280268,0.007857719],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99195427,0.0051395223,0.00034703786,0.0017394881,0.00060442806,0.0002152322],"domain_scores_gemma":[0.9557919,0.038033456,0.0007556354,0.0035214871,0.0012145168,0.0006830218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014422773,0.002297137,0.0015631663,0.0018041974,0.0009654487,0.0026710252,0.0038821467,0.0038916403,0.010007215],"category_scores_gemma":[0.043976113,0.0008952031,0.0022359756,0.0013529132,0.0012218405,0.004137338,0.003582045,0.0060407273,0.0047669285],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0042116996,0.0023547632,0.013298418,0.0036198853,0.0010610258,0.0012987257,0.000987167,0.21256472,0.010270979,0.014492997,0.12648301,0.60935664],"study_design_scores_gemma":[0.0008543658,0.0004815897,0.0027251474,0.00021250648,0.00019776395,0.00059604004,0.00025485244,0.9295089,0.008702934,0.03038345,0.025972541,0.000109990775],"about_ca_topic_score_codex":0.0077163386,"about_ca_topic_score_gemma":0.015222401,"teacher_disagreement_score":0.014422773,"about_ca_system_score_codex":0.00211949,"about_ca_system_score_gemma":0.0036383646,"threshold_uncertainty_score":0.076275826},"labels":[],"label_agreement":null},{"id":"W4401158708","doi":"10.21203/rs.3.rs-4670889/v1","title":"Addressing Gender Bias in Generative Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Generative grammar; Gender bias; Psychology; Cognitive psychology; Social psychology; Computer science; Artificial intelligence","score_opus":0.4433290697922615,"score_gpt":0.48259643743065633,"score_spread":0.039267367638394834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401158708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.097784206,0.0016504469,0.89109063,0.0042170146,0.000305961,0.00007358135,0.0005987627,0.0011063513,0.0031729764],"genre_scores_gemma":[0.89455134,0.0012314615,0.09306154,0.0013706513,0.00092052337,0.00021607826,0.0015050229,0.0013141923,0.0058291527],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913419,0.006741853,0.0002069999,0.00087306567,0.00051341206,0.00032276072],"domain_scores_gemma":[0.86913127,0.1228001,0.0012934001,0.0041421577,0.0018505354,0.00078252394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017960867,0.0011654284,0.0019126385,0.0013642359,0.0013341912,0.0030383551,0.0021541724,0.0029239613,0.005169048],"category_scores_gemma":[0.0982917,0.0015996028,0.0014115215,0.0015077558,0.001569496,0.005573648,0.0033950002,0.0044146483,0.0012536568],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014203951,0.00033892487,0.027122693,0.0007545446,0.00095398707,0.00083265424,0.0034725973,0.31430462,0.0103693465,0.4289597,0.016992275,0.19447835],"study_design_scores_gemma":[0.000056924055,0.000036381884,0.00088861183,0.00003878257,0.00008461333,0.00008959879,0.00011423971,0.8187292,0.0016355744,0.1767896,0.0015154415,0.000020998927],"about_ca_topic_score_codex":0.005067316,"about_ca_topic_score_gemma":0.0076523307,"teacher_disagreement_score":0.017960867,"about_ca_system_score_codex":0.0015349502,"about_ca_system_score_gemma":0.0017546772,"threshold_uncertainty_score":0.09498727},"labels":[],"label_agreement":null},{"id":"W4401340787","doi":"10.1017/s0003055424000716","title":"Improving Probabilistic Models In Text Classification Via Active Learning","year":2024,"lang":"en","type":"article","venue":"American Political Science Review","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Probabilistic logic; Computer science; Artificial intelligence; Machine learning; Active learning (machine learning)","score_opus":0.038489093231234166,"score_gpt":0.328479701192705,"score_spread":0.2899906079614708,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401340787","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005632955,0.0012926232,0.9903882,0.00083497516,0.00012021892,0.00006312302,0.00008591839,0.0007452183,0.00083680346],"genre_scores_gemma":[0.34557372,0.0028652532,0.64252895,0.0012399651,0.0015280038,0.00082523306,0.0012062154,0.00043788215,0.0037947507],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9936249,0.0035131574,0.0003229778,0.0011310875,0.0011695373,0.00023824706],"domain_scores_gemma":[0.95944566,0.03439226,0.0014941987,0.0019901616,0.0023771124,0.00030050616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012418901,0.0019524493,0.00231292,0.0041241255,0.0012930866,0.0040020454,0.0048407833,0.0036730901,0.0021139304],"category_scores_gemma":[0.03643945,0.0012557738,0.0023337468,0.0040240544,0.0022688124,0.00986562,0.0028953888,0.0058446443,0.0017314226],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024476863,0.00044231338,0.0034293411,0.0005089668,0.00032285342,0.000106604944,0.00048397586,0.5134125,0.0024642688,0.052636977,0.0097843185,0.41616324],"study_design_scores_gemma":[0.00001703552,0.000016536622,0.000106009895,0.000019238016,0.00001688207,0.000012936059,0.000010887313,0.9717733,0.0005122487,0.026723273,0.0007810074,0.000010482026],"about_ca_topic_score_codex":0.0030234724,"about_ca_topic_score_gemma":0.0033390454,"teacher_disagreement_score":0.012418901,"about_ca_system_score_codex":0.0016996702,"about_ca_system_score_gemma":0.0012875078,"threshold_uncertainty_score":0.06567818},"labels":[],"label_agreement":null},{"id":"W4401381513","doi":"10.1145/3687273.3687282","title":"Report on the 8th Workshop on Search-Oriented Conversational Artificial Intelligence (SCAI 2024) at CHIIR 2024","year":2024,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Regina; Mila - Quebec Artificial Intelligence Institute","funders":"HORIZON EUROPE Framework Programme; Thüringer Ministerium für Wirtschaft, Wissenschaft und Digitale Gesellschaft; European Commission","keywords":"Computer science; Psychology; Artificial intelligence","score_opus":0.05451247505985401,"score_gpt":0.30534951345901395,"score_spread":0.25083703839915994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401381513","genre_codex":"other","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0388689,0.027057497,0.2197042,0.16536745,0.1396963,0.01023366,0.03657874,0.017520139,0.34497312],"genre_scores_gemma":[0.04546311,0.005346105,0.07654588,0.013424675,0.011861197,0.0046966104,0.04097833,0.0066467007,0.7950374],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9856553,0.0054464648,0.00043669017,0.0016795709,0.0049941097,0.00178779],"domain_scores_gemma":[0.966475,0.0063109794,0.00045973883,0.002752565,0.014653827,0.0093479315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02616235,0.0025148175,0.002020993,0.0023709352,0.0043120584,0.010294052,0.0036156094,0.004401847,0.17093842],"category_scores_gemma":[0.026985202,0.0010718447,0.0023343707,0.0018755783,0.0011865365,0.008987127,0.010634058,0.007123395,0.10192947],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003441457,0.0003343202,0.00046786008,0.00019840692,0.000032837717,0.00011355943,0.00088125665,0.00035233292,0.0023194177,0.0020171348,0.9430107,0.04992802],"study_design_scores_gemma":[0.00012555224,0.0003458488,0.0018685315,0.00021258212,0.000051834566,0.00011681635,0.0013460247,0.0018118537,0.002856854,0.003677643,0.98747665,0.00010992284],"about_ca_topic_score_codex":0.015933217,"about_ca_topic_score_gemma":0.026490046,"teacher_disagreement_score":0.17093842,"about_ca_system_score_codex":0.0033594654,"about_ca_system_score_gemma":0.0077718906,"threshold_uncertainty_score":0.571846},"labels":[],"label_agreement":null},{"id":"W4401408800","doi":"10.1145/3673038.3673124","title":"Arlo: Serving Transformer-based Language Models with Dynamic Input Lengths","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Computer science; Latency (audio); Compiler; Padding; Scheduling (production processes); Parallel computing; Testbed; Distributed computing; Queue; Runtime system; Serialization; Programming language; Computer network; Mathematical optimization","score_opus":0.011367580558465354,"score_gpt":0.2410187164644272,"score_spread":0.22965113590596187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401408800","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059474356,0.00021633985,0.80764025,0.0004942489,0.00010669209,0.000252005,0.001709828,0.12507355,0.0050327033],"genre_scores_gemma":[0.56216204,0.00026146884,0.4204668,0.00044410976,0.00006955232,0.00028311455,0.004346842,0.0064075394,0.005558572],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990355,0.00025419483,0.000074352814,0.00026025614,0.00025392274,0.00012181348],"domain_scores_gemma":[0.99675804,0.0015351671,0.00017935158,0.0010820586,0.0003057919,0.000139622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014637385,0.0009484487,0.0006697579,0.0006090303,0.0004728977,0.001611708,0.0026881427,0.00076343457,0.004592985],"category_scores_gemma":[0.005844004,0.00078299735,0.0010975975,0.00071204285,0.0008046267,0.0034718895,0.0017621691,0.0015534298,0.0022904137],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002188067,0.0010158039,0.010834649,0.001011863,0.00032088548,0.00076530717,0.0015028107,0.48043376,0.102045245,0.05130155,0.052824415,0.29575562],"study_design_scores_gemma":[0.00006348155,0.000059553953,0.00017132354,0.0000056790614,0.000020109459,0.00004827561,0.00006994568,0.97666305,0.011838354,0.0063503473,0.0046910453,0.000018723853],"about_ca_topic_score_codex":0.008936595,"about_ca_topic_score_gemma":0.02252928,"teacher_disagreement_score":0.008936595,"about_ca_system_score_codex":0.0013271424,"about_ca_system_score_gemma":0.0026346406,"threshold_uncertainty_score":0.017769158},"labels":[],"label_agreement":null},{"id":"W4401414318","doi":"10.1109/icra57147.2024.10610065","title":"ISR-LLM: Iterative Self-Refined Large Language Model for Long-Horizon Sequential Task Planning","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":57,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"JST-Mirai Program; Natural Sciences and Engineering Research Council of Canada","keywords":"Task (project management); Computer science; Horizon; Time horizon; Artificial intelligence; Mathematical optimization; Engineering; Mathematics; Systems engineering","score_opus":0.026416638691015636,"score_gpt":0.3070668637932958,"score_spread":0.28065022510228016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401414318","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049168174,0.00011218738,0.9843592,0.00017041125,0.000028480505,0.00016564241,0.00026520118,0.008797391,0.0011846474],"genre_scores_gemma":[0.16184025,0.00014824704,0.8330274,0.00019340408,0.000023876783,0.00061400357,0.0012660341,0.0010214131,0.0018653643],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99845064,0.0006030898,0.00012433337,0.00028277925,0.00042201421,0.00011700587],"domain_scores_gemma":[0.9973067,0.0015735402,0.00020332968,0.0005144763,0.0002961323,0.00010587498],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021770503,0.0012229304,0.00088273,0.0008248173,0.0005868163,0.0011855447,0.0029574924,0.001070879,0.0049988087],"category_scores_gemma":[0.0067225057,0.0008028322,0.0019495302,0.0006224334,0.0012594776,0.0024045098,0.0026312596,0.0023366984,0.0015028664],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033066113,0.00020292815,0.00090769405,0.00049964536,0.00008573397,0.0002893113,0.0007927191,0.7695464,0.010315064,0.03671846,0.009866292,0.17044517],"study_design_scores_gemma":[0.000032728003,0.00004154113,0.000053719727,0.000013791131,0.000012877754,0.000023622222,0.000027062952,0.9842868,0.0022980035,0.009808131,0.0033889913,0.0000127717885],"about_ca_topic_score_codex":0.01355092,"about_ca_topic_score_gemma":0.02352795,"teacher_disagreement_score":0.01355092,"about_ca_system_score_codex":0.0014998621,"about_ca_system_score_gemma":0.004550663,"threshold_uncertainty_score":0.026944041},"labels":[],"label_agreement":null},{"id":"W4401419173","doi":"10.1007/978-3-031-66462-5_3","title":"Sorcerer’s Apprentice? Exploring an AI-Driven Tool to Analyze Academic Texts","year":2024,"lang":"en","type":"book-chapter","venue":"Cognition and exploratory learning in the digital age","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Apprenticeship; Computer science; Mathematics education; Art; History; Psychology; Archaeology","score_opus":0.06850270548865223,"score_gpt":0.2705772373137657,"score_spread":0.20207453182511348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401419173","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19643289,0.002135092,0.656656,0.005989196,0.00052597164,0.00030209657,0.0014386376,0.012283954,0.12423614],"genre_scores_gemma":[0.46321353,0.0011403986,0.4535132,0.00075010885,0.00026094163,0.00020988997,0.0023735603,0.0030149752,0.07552329],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994721,0.00021548646,0.000020750273,0.00014346105,0.00012081253,0.00002744883],"domain_scores_gemma":[0.9970957,0.0021149155,0.00013576583,0.00022987575,0.00021198283,0.00021185956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012671533,0.000511177,0.0002986655,0.002109993,0.0009311599,0.0055230227,0.00073795684,0.00067566347,0.011352928],"category_scores_gemma":[0.006047938,0.00022462744,0.00038179694,0.0015385228,0.0010076228,0.0072626607,0.0015178265,0.0009772899,0.0036969164],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022616121,0.00029139212,0.009040899,0.000443656,0.0000597572,0.0007070373,0.059467442,0.0017543924,0.017249094,0.16119492,0.05102385,0.69854134],"study_design_scores_gemma":[0.00005266632,0.00016194436,0.010577733,0.00041028837,0.00008556108,0.0020099168,0.03806039,0.09385396,0.021013651,0.18256527,0.6511026,0.00010607443],"about_ca_topic_score_codex":0.0009782698,"about_ca_topic_score_gemma":0.002334131,"teacher_disagreement_score":0.011352928,"about_ca_system_score_codex":0.00046166568,"about_ca_system_score_gemma":0.00067761773,"threshold_uncertainty_score":0.037979305},"labels":[],"label_agreement":null},{"id":"W4401499957","doi":"10.5539/elt.v17n9p14","title":"Prompt Engineering for Applied Linguistics: Elements, Examples, Techniques, and Strategies","year":2024,"lang":"en","type":"article","venue":"English Language Teaching","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Applied linguistics; Perspective (graphical); Persona; Linguistics; Psychology; Computer science; Artificial intelligence; Human–computer interaction; Philosophy","score_opus":0.009857271124992097,"score_gpt":0.2530641810576131,"score_spread":0.24320690993262098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401499957","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008893284,0.0031769038,0.9625624,0.007952204,0.00016889401,0.00031472772,0.000048699163,0.0015652973,0.0153176645],"genre_scores_gemma":[0.16730642,0.0034278317,0.82299477,0.00083964656,0.00013226506,0.00057307817,0.00008434206,0.00041227034,0.004229299],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9855377,0.012295083,0.000486967,0.00056135096,0.0009436569,0.00017526382],"domain_scores_gemma":[0.9741697,0.021072123,0.0006604702,0.002498386,0.0011852182,0.00041413892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014158215,0.0013244862,0.00061491516,0.0019989042,0.0018442922,0.006538769,0.0017485322,0.0021034682,0.004144421],"category_scores_gemma":[0.033155948,0.0006007407,0.00064785563,0.001644567,0.00811932,0.010353098,0.0048595225,0.0033856332,0.0017577049],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011014015,0.00017614487,0.0014392451,0.0010632756,0.000019149677,0.00033336674,0.026318967,0.0038554105,0.00428968,0.6841571,0.00688316,0.27135438],"study_design_scores_gemma":[0.00005725363,0.00015157029,0.00042267804,0.0010638902,0.000030842777,0.0008004026,0.009891593,0.034590926,0.007209149,0.7652855,0.1804004,0.00009575404],"about_ca_topic_score_codex":0.0009932043,"about_ca_topic_score_gemma":0.0014613834,"teacher_disagreement_score":0.014158215,"about_ca_system_score_codex":0.0024943585,"about_ca_system_score_gemma":0.0027433308,"threshold_uncertainty_score":0.074876726},"labels":[],"label_agreement":null},{"id":"W4401543468","doi":"10.1145/3643991.3645075","title":"The role of library versions in Developer-ChatGPT conversations","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; World Wide Web; Human–computer interaction","score_opus":0.010239131719209817,"score_gpt":0.21219511470684654,"score_spread":0.20195598298763673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401543468","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4638505,0.0050951485,0.40687835,0.014607016,0.0010274785,0.00047483333,0.0012864635,0.02584197,0.08093828],"genre_scores_gemma":[0.9277651,0.00090853055,0.049093816,0.0017095441,0.00032502096,0.00037343224,0.000896652,0.008280091,0.010647733],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.94187856,0.042906698,0.0024992186,0.0041684303,0.0070128962,0.0015340766],"domain_scores_gemma":[0.68189096,0.24859944,0.017098602,0.029126871,0.016098235,0.007185871],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.045604557,0.0013306628,0.0010571596,0.0059379246,0.0049592983,0.013324164,0.003090998,0.0040210267,0.006548825],"category_scores_gemma":[0.27789,0.0029592589,0.0008204037,0.0030899872,0.0045788516,0.030336268,0.012213925,0.004517482,0.0032151132],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018592615,0.00035406323,0.08940477,0.0014728301,0.00024865486,0.004207535,0.36146745,0.003799347,0.018354148,0.090251446,0.028675485,0.399905],"study_design_scores_gemma":[0.00041509038,0.0010101479,0.10153864,0.004172698,0.0007922921,0.007888271,0.11636419,0.11541515,0.030995244,0.1721875,0.44756234,0.0016584587],"about_ca_topic_score_codex":0.0060411254,"about_ca_topic_score_gemma":0.0048031267,"teacher_disagreement_score":0.045604557,"about_ca_system_score_codex":0.004832477,"about_ca_system_score_gemma":0.0037147913,"threshold_uncertainty_score":0.2411828},"labels":[],"label_agreement":null},{"id":"W4401567109","doi":"10.1177/09567976241254037","title":"The Language of (Non)Replicable Social Science","year":2024,"lang":"en","type":"article","venue":"Psychological Science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Persuasion; Psychology; Set (abstract data type); Replication (statistics); Social psychology; Cognitive psychology; Linguistics; Computer science","score_opus":0.04289174541375712,"score_gpt":0.38909850360854553,"score_spread":0.3462067581947884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401567109","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48042452,0.011340483,0.4567352,0.008650103,0.002219514,0.011518619,0.010291553,0.001133553,0.017686374],"genre_scores_gemma":[0.8834119,0.000822816,0.097174324,0.0012295372,0.00047862148,0.013530038,0.0022100147,0.00024400899,0.0008987713],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.26166663,0.5879682,0.0903781,0.022580085,0.03550015,0.001906916],"domain_scores_gemma":[0.04176549,0.71030915,0.07281338,0.15229395,0.022082573,0.00073552574],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.54742503,0.0011614165,0.002002383,0.01173064,0.0029598651,0.012653393,0.004246744,0.0034259423,0.003896943],"category_scores_gemma":[0.8138274,0.0013742849,0.002463866,0.0133378,0.015918301,0.0070886244,0.0061569763,0.0041688583,0.0011597471],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004955741,0.0010019172,0.19907175,0.02756421,0.005495179,0.002924441,0.27372786,0.0049184156,0.023675213,0.18884754,0.0145118395,0.25330594],"study_design_scores_gemma":[0.0019668653,0.003874343,0.2370078,0.016586993,0.003771853,0.005395765,0.053114783,0.028292838,0.043679737,0.45476353,0.15002298,0.0015225493],"about_ca_topic_score_codex":0.0009897918,"about_ca_topic_score_gemma":0.001098584,"teacher_disagreement_score":0.45257497,"about_ca_system_score_codex":0.0033721512,"about_ca_system_score_gemma":0.0051551727,"threshold_uncertainty_score":0.5581056},"labels":[],"label_agreement":null},{"id":"W4401587116","doi":"10.20944/preprints202408.0334.v1","title":"Promptology: Enhancing Human-AI Interaction in Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Natural language processing; Linguistics; Artificial intelligence; Philosophy","score_opus":0.10697286213527459,"score_gpt":0.387377614746558,"score_spread":0.28040475261128345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401587116","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06352777,0.0005034269,0.90108377,0.0019213933,0.00013471818,0.00043440968,0.00038029224,0.008174064,0.023840185],"genre_scores_gemma":[0.5048762,0.00067185395,0.47633737,0.00041642433,0.00007056743,0.0006214283,0.0009721821,0.00166357,0.014370431],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99743384,0.0017591828,0.000081637976,0.00032376204,0.00031066188,0.00009082664],"domain_scores_gemma":[0.9906833,0.006926644,0.000254926,0.0014553623,0.00034937458,0.00033038206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043498124,0.000740247,0.0004807958,0.0010452131,0.0011554362,0.003954365,0.0012655008,0.0013772714,0.013069038],"category_scores_gemma":[0.016746698,0.0004958769,0.00085513806,0.0007092509,0.002391293,0.008647558,0.007003401,0.0018433636,0.0025424226],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007406842,0.0008437323,0.008704371,0.0019809096,0.00013853642,0.0012036798,0.05366484,0.06288511,0.04309938,0.3407129,0.020988325,0.46503752],"study_design_scores_gemma":[0.00020669377,0.0006440061,0.00223198,0.00040675336,0.00012040462,0.0010033996,0.009570927,0.43966419,0.028452018,0.29995182,0.21759827,0.00014943597],"about_ca_topic_score_codex":0.0010871668,"about_ca_topic_score_gemma":0.0023372308,"teacher_disagreement_score":0.013069038,"about_ca_system_score_codex":0.0010823631,"about_ca_system_score_gemma":0.0014765636,"threshold_uncertainty_score":0.043720245},"labels":[],"label_agreement":null},{"id":"W4401597819","doi":"10.1371/journal.pone.0307741","title":"GPT-4 as an X data annotator: Unraveling its performance on a stance classification task","year":2024,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Lakehead University","funders":"Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Natural language processing; Task (project management); Annotation; Machine learning; Context (archaeology); Benchmark (surveying); Generalizability theory; Set (abstract data type); Psychology","score_opus":0.19554144764206105,"score_gpt":0.30028207906352067,"score_spread":0.10474063142145962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401597819","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43873188,0.0038176721,0.4307419,0.0027913847,0.0027425317,0.0015507019,0.0054566,0.08962689,0.024540443],"genre_scores_gemma":[0.54959905,0.0005775729,0.41533822,0.0018231557,0.00027097674,0.0012693591,0.013767854,0.0030794886,0.014274381],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9932401,0.0030269003,0.00032905565,0.0021467581,0.00084668165,0.00041046398],"domain_scores_gemma":[0.9840307,0.008427859,0.00058203243,0.0033784818,0.0028685452,0.00071225915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010398374,0.002569637,0.0014167898,0.001574035,0.0016529197,0.0026012533,0.002968923,0.0032306411,0.0051821955],"category_scores_gemma":[0.028050402,0.0006797156,0.0012972826,0.0011631493,0.0013566389,0.0050029703,0.005107576,0.0042507527,0.0070942147],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040769503,0.0010401628,0.021007894,0.002098484,0.00052151503,0.0012169279,0.0047368514,0.04014061,0.063038945,0.004840044,0.07873989,0.77854174],"study_design_scores_gemma":[0.00050144794,0.0014389328,0.009301189,0.00036069044,0.00024610947,0.0009464209,0.0027587058,0.8728382,0.059906386,0.011199801,0.040230982,0.0002711302],"about_ca_topic_score_codex":0.008487102,"about_ca_topic_score_gemma":0.01351957,"teacher_disagreement_score":0.010398374,"about_ca_system_score_codex":0.0014890813,"about_ca_system_score_gemma":0.0023795813,"threshold_uncertainty_score":0.054992497},"labels":[],"label_agreement":null},{"id":"W4401752756","doi":"10.1109/ichi61247.2024.00089","title":"Seeing Beyond Borders: Evaluating LLMs in Multilingual Ophthalmological Question Answering","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Question answering; Computer science; Natural language processing; Artificial intelligence","score_opus":0.0550715705987934,"score_gpt":0.39265617320026674,"score_spread":0.3375846026014733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401752756","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7970575,0.0139125995,0.12733598,0.0033932892,0.0010873255,0.0024966286,0.014660465,0.022480285,0.017575784],"genre_scores_gemma":[0.86590856,0.0009420548,0.10212708,0.0009876916,0.00033389745,0.00088319456,0.026649546,0.0003261898,0.0018417679],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9914892,0.0049304008,0.0008285678,0.0016528032,0.0008322559,0.00026693876],"domain_scores_gemma":[0.9666021,0.028969105,0.0007238337,0.001160007,0.0014661994,0.0010788202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009937399,0.0018777703,0.0009147196,0.002721221,0.0008969547,0.0033999425,0.0018415082,0.003566447,0.0063820705],"category_scores_gemma":[0.04475634,0.0003409513,0.0013965832,0.0013974959,0.0007918563,0.0059463712,0.0045524365,0.0022165047,0.0029458168],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012575602,0.00429309,0.04818094,0.0067005246,0.0019974473,0.0010452494,0.006390338,0.08669194,0.016638706,0.0042375396,0.05182247,0.7594261],"study_design_scores_gemma":[0.0012590832,0.0040148203,0.025176693,0.0005702294,0.0009407457,0.0008416344,0.005976605,0.90711904,0.01241116,0.016768131,0.024660865,0.00026099934],"about_ca_topic_score_codex":0.009454943,"about_ca_topic_score_gemma":0.010483069,"teacher_disagreement_score":0.009937399,"about_ca_system_score_codex":0.0016811289,"about_ca_system_score_gemma":0.0017996964,"threshold_uncertainty_score":0.052554667},"labels":[],"label_agreement":null},{"id":"W4401768173","doi":"10.1007/978-3-031-66705-3_14","title":"Investigating a Semantic Similarity Loss Function for the Parallel Training of Abstractive and Extractive Scientific Document Summarizers","year":2024,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Information retrieval; Similarity (geometry); Function (biology); Natural language processing; Semantic similarity; Artificial intelligence; Biology; Genetics","score_opus":0.08608398480217001,"score_gpt":0.314668130721619,"score_spread":0.22858414591944898,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401768173","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15172298,0.0015563839,0.83978933,0.0008302472,0.00021102796,0.00013615104,0.00021130426,0.0024122533,0.0031301945],"genre_scores_gemma":[0.6789633,0.0005423392,0.30578387,0.00040840407,0.00031025152,0.00019650835,0.0015632603,0.00044748152,0.011784611],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896014,0.00037217344,0.00008766417,0.00024242558,0.00023646136,0.00010111027],"domain_scores_gemma":[0.9961041,0.002487808,0.0001850176,0.0003073901,0.0007783814,0.00013738069],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044508297,0.001006506,0.0010849657,0.0007152512,0.00043013805,0.001621647,0.0016749382,0.0019781704,0.0023743866],"category_scores_gemma":[0.008002639,0.00044872603,0.00076573616,0.00069570274,0.00042879398,0.0030796796,0.0013625122,0.0021288393,0.00098849],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001245754,0.00074184465,0.003391481,0.0002953357,0.00025898838,0.00014041047,0.00016967226,0.22999662,0.034917,0.0073915333,0.0063106306,0.71514076],"study_design_scores_gemma":[0.000024240384,0.00021035048,0.00049369456,0.000008480537,0.000035507004,0.000031942975,0.00003348623,0.9893875,0.007831746,0.0014261233,0.00050934934,0.000007575727],"about_ca_topic_score_codex":0.0026665463,"about_ca_topic_score_gemma":0.0030813299,"teacher_disagreement_score":0.0044508297,"about_ca_system_score_codex":0.0007317366,"about_ca_system_score_gemma":0.0012620414,"threshold_uncertainty_score":0.02353853},"labels":[],"label_agreement":null},{"id":"W4401783591","doi":"10.1080/03155986.2024.2388452","title":"LM4OPT: Unveiling the potential of Large Language Models in formulating mathematical optimization problems","year":2024,"lang":"en","type":"article","venue":"INFOR Information Systems and Operational Research","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Comprehension; Task (project management); Computer science; Shot (pellet); Natural language processing; Artificial intelligence; Machine learning; Engineering; Chemistry","score_opus":0.0502706093629151,"score_gpt":0.34489316682392085,"score_spread":0.29462255746100574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401783591","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10499034,0.0024952118,0.8443766,0.0021030845,0.00046012423,0.00031776787,0.0013361913,0.034586214,0.009334491],"genre_scores_gemma":[0.46365348,0.0006016399,0.5223021,0.001530127,0.00014149505,0.0005652437,0.0037182106,0.0026218176,0.0048658648],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988391,0.0005077556,0.00006637806,0.0003264947,0.0001830879,0.00007725668],"domain_scores_gemma":[0.9975769,0.001591584,0.00009573384,0.00045115783,0.00019864096,0.00008601436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023366439,0.002297432,0.0010976938,0.0007366138,0.00050845183,0.0020875263,0.0030115542,0.0021068116,0.0042339712],"category_scores_gemma":[0.011202972,0.0008917837,0.001511112,0.00067346194,0.0011449209,0.004133972,0.002696943,0.0040739705,0.0020663855],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035054318,0.0002747052,0.0019254723,0.0006371443,0.00018692516,0.00033748755,0.00023916073,0.75455064,0.005833295,0.013410388,0.020017657,0.2022366],"study_design_scores_gemma":[0.000029304078,0.0000553537,0.00007693849,0.000014723525,0.000012157914,0.000031287702,0.000026006035,0.99019825,0.0013092879,0.00623797,0.001998634,0.000010036369],"about_ca_topic_score_codex":0.0068499027,"about_ca_topic_score_gemma":0.014194148,"teacher_disagreement_score":0.0068499027,"about_ca_system_score_codex":0.0013362634,"about_ca_system_score_gemma":0.0022077826,"threshold_uncertainty_score":0.01416409},"labels":[],"label_agreement":null},{"id":"W4401829805","doi":"10.18280/ria.380421","title":"Enhancing Question Generation in Bahasa Using Pretrained Language Models","year":2024,"lang":"en","type":"article","venue":"Revue d intelligence artificielle","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Binus University","keywords":"Computer science; Natural language processing; Linguistics; Language model; Artificial intelligence; Psychology; Philosophy","score_opus":0.07074700845449783,"score_gpt":0.3116321930750804,"score_spread":0.24088518462058256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401829805","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5420655,0.0042025857,0.34862483,0.0024352754,0.000735581,0.0014091632,0.011614543,0.07191168,0.017000964],"genre_scores_gemma":[0.822943,0.0005461356,0.15008359,0.00067354616,0.00008246452,0.00045926732,0.01984532,0.0005473552,0.0048193103],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9987558,0.0005680463,0.00009193652,0.00039644018,0.00011987125,0.00006784277],"domain_scores_gemma":[0.9946955,0.0038444991,0.00016586296,0.0005181569,0.00060633075,0.000169738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028824578,0.0009711268,0.00075366226,0.0009153813,0.0004023268,0.0016005824,0.0013864513,0.0011728927,0.0033522216],"category_scores_gemma":[0.009596926,0.00034319627,0.0010615714,0.0006083434,0.0004183432,0.0033625562,0.0010327784,0.0018338583,0.002798206],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015587708,0.0015853319,0.022676677,0.001985242,0.00029554343,0.0006936692,0.0022764208,0.14736547,0.034299303,0.0050296527,0.038322005,0.7439119],"study_design_scores_gemma":[0.00014175657,0.000458815,0.004738028,0.000071961746,0.00010449315,0.0002806783,0.00049199094,0.9528422,0.021965861,0.0050337105,0.013814495,0.000056033015],"about_ca_topic_score_codex":0.009432635,"about_ca_topic_score_gemma":0.008460319,"teacher_disagreement_score":0.009432635,"about_ca_system_score_codex":0.0013905367,"about_ca_system_score_gemma":0.0013402629,"threshold_uncertainty_score":0.018755496},"labels":[],"label_agreement":null},{"id":"W4401857430","doi":"10.1145/3637528.3671474","title":"A Review of Modern Recommender Systems Using Generative Models (Gen-RecSys)","year":2024,"lang":"en","type":"review","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":94,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Recommender system; Computer science; Generative grammar; Generative model; Artificial intelligence; Key (lock); Multidisciplinary approach; Machine learning; Data science; Information retrieval","score_opus":0.3534450477250533,"score_gpt":0.40607544394428235,"score_spread":0.05263039621922905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401857430","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009591589,0.96564794,0.026066076,0.0015385242,0.0004679279,0.00004321058,0.00027917355,0.00023123191,0.0047667148],"genre_scores_gemma":[0.011004711,0.9608331,0.022710172,0.0009451344,0.0012395929,0.00006912638,0.0006757671,0.000075344244,0.002447045],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990108,0.00031470341,0.00010628815,0.00022100234,0.0002944803,0.000052731288],"domain_scores_gemma":[0.99533445,0.0033317276,0.00014431732,0.00032128065,0.00078219624,0.0000860907],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002703366,0.0013718879,0.001607609,0.0024058188,0.0004529054,0.0016005421,0.0019219476,0.0014786477,0.0053925947],"category_scores_gemma":[0.007163253,0.0010144298,0.0014135783,0.0053877095,0.00048674672,0.0030692269,0.001103106,0.0016727566,0.005250554],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006168792,0.00010588718,0.0012962142,0.009642384,0.00033093043,0.000088268105,0.0001452098,0.006710983,0.0007202828,0.02123026,0.046867535,0.9128004],"study_design_scores_gemma":[0.000035210334,0.00033148302,0.0024525807,0.0047475896,0.00056466943,0.001125932,0.00017124286,0.01839652,0.0011572333,0.025432497,0.9454442,0.00014083367],"about_ca_topic_score_codex":0.0055146795,"about_ca_topic_score_gemma":0.0056992862,"teacher_disagreement_score":0.0055146795,"about_ca_system_score_codex":0.0010137019,"about_ca_system_score_gemma":0.0019527384,"threshold_uncertainty_score":0.018040001},"labels":[],"label_agreement":null},{"id":"W4401863228","doi":"10.1145/3637528.3671793","title":"<scp>Binder:</scp> Hierarchical Concept Representation through Order Embedding of Binary Vectors","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Embedding; Representation (politics); Theoretical computer science; Computer science; Simple (philosophy); Mathematics; Algorithm; Artificial intelligence","score_opus":0.03332678432754565,"score_gpt":0.3208398496670251,"score_spread":0.28751306533947946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401863228","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027787115,0.0011005566,0.8791311,0.0032763933,0.0012711076,0.0004259336,0.015737684,0.06615048,0.03012808],"genre_scores_gemma":[0.078197695,0.0024523798,0.7533273,0.002750687,0.0009185337,0.0014454317,0.07686068,0.016797336,0.06724996],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99792457,0.00033246417,0.00012690578,0.00041467874,0.0010741956,0.00012713157],"domain_scores_gemma":[0.99604416,0.0009041402,0.00020473194,0.0016189029,0.0010097692,0.00021841687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016788385,0.002542447,0.0012306811,0.0031553977,0.001062522,0.0041044857,0.0047583487,0.0022642629,0.11031891],"category_scores_gemma":[0.010479603,0.00094004377,0.0015845131,0.005213397,0.0015752052,0.009553378,0.0054456065,0.0038030015,0.061495457],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001474933,0.00009351634,0.0002506255,0.00047652714,0.0000594141,0.00023203599,0.00018829612,0.0062592574,0.0064787725,0.053095262,0.6137277,0.31899107],"study_design_scores_gemma":[0.000107465996,0.00010013771,0.0007261282,0.00018238806,0.0000428484,0.00059403596,0.0001216266,0.32315248,0.025250696,0.17711954,0.47245315,0.00014942048],"about_ca_topic_score_codex":0.012567983,"about_ca_topic_score_gemma":0.016753545,"teacher_disagreement_score":0.11031891,"about_ca_system_score_codex":0.0016428126,"about_ca_system_score_gemma":0.0020722374,"threshold_uncertainty_score":0.36905348},"labels":[],"label_agreement":null},{"id":"W4401869038","doi":"10.1111/exsy.13712","title":"Intent detection for task‐oriented conversational agents: A comparative study of recurrent neural networks and transformer models","year":2024,"lang":"en","type":"article","venue":"Expert Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Computer science; Transformer; Task (project management); Artificial neural network; Artificial intelligence; Machine learning; Natural language processing; Voltage; Systems engineering","score_opus":0.07955754620462314,"score_gpt":0.3157993161297343,"score_spread":0.23624176992511114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401869038","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21875365,0.027157776,0.73898095,0.0012321798,0.00029291018,0.00013413679,0.00023539875,0.0017579808,0.011454988],"genre_scores_gemma":[0.9456752,0.005330465,0.04646734,0.00014254707,0.00009389692,0.000050291907,0.0002768029,0.000050961946,0.0019123794],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99943954,0.00021634717,0.000040855204,0.00014746322,0.00010976348,0.00004609175],"domain_scores_gemma":[0.99776566,0.001496569,0.00012775374,0.00009266908,0.00045702056,0.000060401755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019204069,0.0008053277,0.00055323535,0.00068267726,0.00018026985,0.0010872739,0.00089145114,0.0006583086,0.00083951675],"category_scores_gemma":[0.0041948175,0.0002546204,0.00060396,0.00039699202,0.00030700053,0.0013536252,0.00046427123,0.00090328837,0.00037908377],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006092717,0.0003755413,0.009896437,0.00088337297,0.0006032322,0.00027285356,0.0005152907,0.39167935,0.010711035,0.016856939,0.0032245426,0.5643722],"study_design_scores_gemma":[0.0000036285307,0.000090493544,0.0007155346,0.000030479494,0.00005523989,0.000028724568,0.000037929196,0.9950145,0.001301868,0.0020759427,0.00063492625,0.000010715616],"about_ca_topic_score_codex":0.006083136,"about_ca_topic_score_gemma":0.00401329,"teacher_disagreement_score":0.006083136,"about_ca_system_score_codex":0.00068452064,"about_ca_system_score_gemma":0.00049531827,"threshold_uncertainty_score":0.012095451},"labels":[],"label_agreement":null},{"id":"W4401943459","doi":"10.1109/icdh62654.2024.00030","title":"Comparative Analysis of Open-Source Language Models in Summarizing Medical Text Data","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Open source; Natural language processing; Data science; Data modeling; Information retrieval; Programming language; Database; Software","score_opus":0.11807377448352604,"score_gpt":0.38186888981358724,"score_spread":0.26379511533006117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401943459","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6962438,0.015739229,0.20945266,0.005658106,0.0011735115,0.0015617906,0.015961511,0.040101077,0.014108303],"genre_scores_gemma":[0.7839047,0.0029032393,0.16455166,0.0006715051,0.00040009545,0.0012002217,0.04159307,0.0020602252,0.0027152498],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9819266,0.011503308,0.0019351521,0.0016663411,0.0025345804,0.0004340229],"domain_scores_gemma":[0.8614308,0.11528607,0.0031192808,0.006567315,0.012090405,0.0015061289],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026835477,0.0020904983,0.001020812,0.0070970403,0.0010352633,0.0032889254,0.002479598,0.0021542863,0.0018095389],"category_scores_gemma":[0.0883228,0.00056218856,0.0023282773,0.003678453,0.0008794165,0.006091069,0.0031450298,0.0020269754,0.001243744],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009261334,0.003734215,0.04743353,0.0068511036,0.0041900226,0.00069946283,0.004997456,0.20453957,0.015096927,0.0097805,0.0399093,0.65350664],"study_design_scores_gemma":[0.00062703283,0.0025885093,0.025378224,0.00059228437,0.0013469007,0.00032132302,0.0022375719,0.91963166,0.016332375,0.012281978,0.01832508,0.00033706415],"about_ca_topic_score_codex":0.01107334,"about_ca_topic_score_gemma":0.01295958,"teacher_disagreement_score":0.026835477,"about_ca_system_score_codex":0.0025098252,"about_ca_system_score_gemma":0.0026380408,"threshold_uncertainty_score":0.14192122},"labels":[],"label_agreement":null},{"id":"W4401943674","doi":"10.1109/icdh62654.2024.00031","title":"Enhancing Large Language Models with Human Expertise for Disease Detection in Electronic Health Records","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Calgary","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Health records; Electronic health record; Human disease; Disease; Data science; Natural language processing; Health care; Medicine; Pathology","score_opus":0.015011222716406872,"score_gpt":0.2914003633171336,"score_spread":0.2763891406007267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401943674","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.071373574,0.0006821779,0.9151322,0.0015793897,0.000071885726,0.00027049432,0.0011666755,0.008612071,0.0011115605],"genre_scores_gemma":[0.55723363,0.00044568078,0.4333064,0.0010320118,0.0002226595,0.00042102707,0.0047348477,0.0005615545,0.002042164],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9966415,0.0019873746,0.00017669234,0.0007758716,0.0002688156,0.00014973035],"domain_scores_gemma":[0.9696444,0.027601693,0.0005810566,0.0008884173,0.0009979384,0.00028650265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007063882,0.001397291,0.00084245246,0.0021686864,0.0007744602,0.0020041848,0.0016557107,0.0013499664,0.0018076554],"category_scores_gemma":[0.024340322,0.00082130794,0.0023248086,0.0010629259,0.0006074469,0.0027226575,0.0023803215,0.0026510493,0.0016807561],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096663064,0.0011968126,0.0307157,0.0007913307,0.00082848774,0.00083694607,0.0033770644,0.35450962,0.016699366,0.009856369,0.020308428,0.5599132],"study_design_scores_gemma":[0.000027723978,0.00003980114,0.00078144495,0.000020232948,0.00005949443,0.000049201924,0.00008388539,0.98959255,0.001450619,0.0068559684,0.0010207449,0.000018323013],"about_ca_topic_score_codex":0.014652457,"about_ca_topic_score_gemma":0.026307572,"teacher_disagreement_score":0.014652457,"about_ca_system_score_codex":0.001311137,"about_ca_system_score_gemma":0.0021102661,"threshold_uncertainty_score":0.037357807},"labels":[],"label_agreement":null},{"id":"W4402008279","doi":"10.31235/osf.io/wg82k","title":"Updating “The Future of Coding”: Qualitative Coding with Generative Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Generative grammar; Coding (social sciences); Computer science; Natural language processing; Generative model; Artificial intelligence; Linguistics; Mathematics; Statistics; Philosophy","score_opus":0.04143052169491773,"score_gpt":0.3302597114237466,"score_spread":0.2888291897288289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402008279","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009987927,0.00038463247,0.97185206,0.008509685,0.0002898029,0.00021391948,0.0006269579,0.0009262388,0.007208823],"genre_scores_gemma":[0.3044109,0.0005103986,0.68586314,0.0025430415,0.00018785059,0.0012719185,0.0009240405,0.00142165,0.002867078],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.91457486,0.07312189,0.0017363367,0.0046059685,0.005267912,0.0006930139],"domain_scores_gemma":[0.6655261,0.2557914,0.008587294,0.050523285,0.017836235,0.0017358088],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06481946,0.0012478759,0.00075558457,0.0031831658,0.0032421858,0.009975852,0.003967748,0.0024250285,0.009984042],"category_scores_gemma":[0.3187526,0.0013367662,0.0014771557,0.0031468289,0.017723626,0.021271681,0.0080855535,0.0058695357,0.0027022169],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015727944,0.000053418633,0.0039861053,0.00072929746,0.00007290308,0.00017347098,0.050710186,0.009883768,0.001830337,0.82570255,0.010740333,0.09596028],"study_design_scores_gemma":[0.00004648574,0.000026754664,0.00054893235,0.0005809358,0.000024381223,0.00013766401,0.0057614837,0.046189863,0.0027552138,0.90042734,0.043403644,0.00009730907],"about_ca_topic_score_codex":0.0057914956,"about_ca_topic_score_gemma":0.006537707,"teacher_disagreement_score":0.93518054,"about_ca_system_score_codex":0.006824589,"about_ca_system_score_gemma":0.008389481,"threshold_uncertainty_score":0.34280217},"labels":[],"label_agreement":null},{"id":"W4402081416","doi":"10.1007/978-3-031-70563-2_10","title":"Analyzing Biases in Popular Answer Selection Datasets on Neural-Based QA Models","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Selection (genetic algorithm); Artificial intelligence; Machine learning; Artificial neural network; Natural language processing","score_opus":0.044718066341635436,"score_gpt":0.2792266015176221,"score_spread":0.23450853517598663,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402081416","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.960194,0.0038426202,0.018311387,0.0015046926,0.00023637878,0.000100890065,0.0111846775,0.001333796,0.0032916653],"genre_scores_gemma":[0.9638591,0.00038344832,0.009891319,0.00024927885,0.00015313707,0.000083220046,0.0235007,0.00016896272,0.0017107454],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9947488,0.003201865,0.0003356417,0.0006830801,0.0007980915,0.00023256794],"domain_scores_gemma":[0.930404,0.05882964,0.0016606478,0.004300749,0.004187864,0.0006170205],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011654527,0.00069385866,0.0008511336,0.002140737,0.0007148873,0.0015707151,0.0013219089,0.0016315231,0.0027706851],"category_scores_gemma":[0.0589851,0.00029453394,0.00085653784,0.00249227,0.0006724803,0.002750251,0.0012828802,0.0018748555,0.0012636731],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011198372,0.002361983,0.2389833,0.0027888392,0.0019969281,0.0005033033,0.0023179667,0.18864284,0.01447738,0.017580068,0.14337277,0.3757763],"study_design_scores_gemma":[0.0003466256,0.0005794082,0.049929623,0.00022697925,0.00032218813,0.00028341715,0.00067183364,0.9118737,0.0073059956,0.019355131,0.009024906,0.00008021158],"about_ca_topic_score_codex":0.0056727235,"about_ca_topic_score_gemma":0.00857961,"teacher_disagreement_score":0.98834544,"about_ca_system_score_codex":0.0015257855,"about_ca_system_score_gemma":0.0007543629,"threshold_uncertainty_score":0.061635733},"labels":[],"label_agreement":null},{"id":"W4402124327","doi":"10.1109/access.2024.3453215","title":"A Comprehensive Evaluation of Neural SPARQL Query Generation From Natural Language Questions","year":2024,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; SPARQL; RDF query language; Natural language generation; Query language; Natural language; Information retrieval; Natural language processing; Web search query; Artificial intelligence; Natural language user interface; Web query classification; Search engine; Semantic Web; RDF","score_opus":0.11068454227856032,"score_gpt":0.3852054805728277,"score_spread":0.2745209382942674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402124327","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76690346,0.009876964,0.13476802,0.0025013513,0.00068489,0.0016198559,0.010293843,0.050208166,0.02314342],"genre_scores_gemma":[0.82213485,0.002296069,0.13759166,0.000835324,0.00011386713,0.0005116828,0.029131088,0.0014592684,0.005926137],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99548435,0.002065586,0.0003942871,0.00078693684,0.0011046766,0.0001640448],"domain_scores_gemma":[0.9881896,0.008484288,0.00024207421,0.0012214254,0.001630238,0.00023235532],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071418085,0.0013786862,0.0009814958,0.0010366332,0.00067099987,0.0011608879,0.0021331932,0.0014497574,0.004951927],"category_scores_gemma":[0.024889294,0.0004412359,0.0007560898,0.0013896154,0.00064175884,0.002849247,0.0013383911,0.0015753658,0.0018083672],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029947702,0.0023694404,0.009600106,0.0046338374,0.0008602295,0.0005130719,0.00086023804,0.250101,0.023370834,0.0043379483,0.04100829,0.6593503],"study_design_scores_gemma":[0.00035064374,0.001312269,0.004922482,0.0001571342,0.0001722516,0.0003145696,0.00041786468,0.948577,0.027302604,0.0032628437,0.01314113,0.00006927332],"about_ca_topic_score_codex":0.016536472,"about_ca_topic_score_gemma":0.01544875,"teacher_disagreement_score":0.016536472,"about_ca_system_score_codex":0.0020308392,"about_ca_system_score_gemma":0.0017845538,"threshold_uncertainty_score":0.037769973},"labels":[],"label_agreement":null},{"id":"W4402192689","doi":"10.7202/1112894ar","title":"Maîtriser le Chat (ro)botté ou comment soumettre l’intelligence artificielle au service de nos usagers en milieu universitaire ?","year":2024,"lang":"fr","type":"article","venue":"Documentation et bibliothèques","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Université de Montréal; Bibliothèque et Archives nationales du Québec","funders":"","keywords":"Humanities; Geography; Forestry; Art","score_opus":0.05247826218939822,"score_gpt":0.3276500091657381,"score_spread":0.2751717469763399,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402192689","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17110205,0.0067783995,0.46494853,0.023316728,0.00212314,0.00037892963,0.00029560743,0.0056769215,0.32537976],"genre_scores_gemma":[0.7712209,0.0024848797,0.10118092,0.002144051,0.0004800148,0.0002716945,0.00023073483,0.0008234329,0.12116326],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9957562,0.0024251405,0.00015662542,0.0005569688,0.00082797423,0.0002771822],"domain_scores_gemma":[0.9930917,0.0033238588,0.00057987403,0.0011510768,0.001226608,0.0006268551],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0044348496,0.00074844726,0.0005280383,0.0012311231,0.002821035,0.007846655,0.001243837,0.002876794,0.012246684],"category_scores_gemma":[0.014974969,0.00048112636,0.0005277812,0.0009627692,0.0037824863,0.010985958,0.003572856,0.0022062883,0.0057971072],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049551815,0.00018711318,0.010681868,0.0013405785,0.00008381198,0.0014915658,0.17395495,0.0033424136,0.031418107,0.48516098,0.021933023,0.2699101],"study_design_scores_gemma":[0.00006340187,0.00037915228,0.006559091,0.0012103356,0.00012306643,0.002235093,0.06599984,0.025291309,0.014878115,0.081879295,0.8011142,0.0002671101],"about_ca_topic_score_codex":0.0074135745,"about_ca_topic_score_gemma":0.008303412,"teacher_disagreement_score":0.99215335,"about_ca_system_score_codex":0.0024362504,"about_ca_system_score_gemma":0.0020220522,"threshold_uncertainty_score":0.040969312},"labels":[],"label_agreement":null},{"id":"W4402206897","doi":"10.1016/j.nlp.2024.100102","title":"Job description parsing with explainable transformer based ensemble models to extract the technical and non-technical skills","year":2024,"lang":"en","type":"article","venue":"Natural Language Processing Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Interpretability; Computer science; Margin (machine learning); Artificial intelligence; Machine learning; Transformer; Statistical model; Engineering","score_opus":0.013537905842925406,"score_gpt":0.2620085156993179,"score_spread":0.24847060985639247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402206897","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17916663,0.001530385,0.78725547,0.0009602274,0.00028174793,0.00022526618,0.0074997237,0.014512393,0.0085681435],"genre_scores_gemma":[0.7903439,0.0007190886,0.17756374,0.000350965,0.00010028137,0.00023278281,0.018682512,0.0003752851,0.011631552],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997956,0.000038692095,0.0000127286585,0.000085825704,0.00003635396,0.00003075667],"domain_scores_gemma":[0.99951994,0.0002477315,0.000035261746,0.000068779176,0.00010248379,0.000025838868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00050892675,0.0009084745,0.0003600557,0.0012215595,0.00029785736,0.00055249763,0.0009659869,0.0006872932,0.0023853362],"category_scores_gemma":[0.0015454487,0.0003014377,0.0011640537,0.0010430173,0.00021247266,0.0014602509,0.0008710432,0.0013523522,0.0012571058],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042964742,0.00044437178,0.01988093,0.00034343058,0.0002747146,0.000777702,0.00061517453,0.2733017,0.016682891,0.009728899,0.030535165,0.64698535],"study_design_scores_gemma":[0.000009581355,0.00003874311,0.0020299163,0.000021152355,0.000056601395,0.00007824958,0.00005947211,0.98555076,0.0035052174,0.0047537256,0.003881671,0.000014801224],"about_ca_topic_score_codex":0.012681339,"about_ca_topic_score_gemma":0.02418,"teacher_disagreement_score":0.012681339,"about_ca_system_score_codex":0.00063705037,"about_ca_system_score_gemma":0.0009670905,"threshold_uncertainty_score":0.02521503},"labels":[],"label_agreement":null},{"id":"W4402264726","doi":"10.23919/acc60939.2024.10644407","title":"Domain-adaptation with knowledge accumulation through parallel stacked autoencoders: methodology and application to sulfur recovery","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Natural Science Foundation of China","keywords":"Domain adaptation; Computer science; Adaptation (eye); Domain (mathematical analysis); Artificial intelligence; Computer architecture; Parallel computing; Neuroscience; Psychology","score_opus":0.12615666617115254,"score_gpt":0.3633077928187106,"score_spread":0.23715112664755808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402264726","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027699746,0.0002807085,0.97037584,0.000097282646,0.000028104938,0.000029911473,0.000037309732,0.00051261,0.0009383692],"genre_scores_gemma":[0.79549503,0.0004222467,0.20097584,0.00009592993,0.000048268947,0.00012373335,0.00017666609,0.00007559486,0.002586718],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998273,0.000039993214,0.000011386713,0.000059963288,0.00003607154,0.000025144796],"domain_scores_gemma":[0.9995395,0.00022580326,0.00005273917,0.00005460501,0.000109633846,0.000017671544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00072923466,0.0007192265,0.0005359294,0.00038761337,0.00023818127,0.0004421148,0.0008073166,0.0006283092,0.000835884],"category_scores_gemma":[0.0012399182,0.00037286582,0.0007013598,0.00049488054,0.0004738193,0.0008183581,0.000812825,0.001032657,0.00027584107],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005786268,0.0000630158,0.0007986177,0.000053418982,0.00007244965,0.00007934153,0.00007612451,0.84890825,0.008711034,0.0026372075,0.00049698085,0.13804567],"study_design_scores_gemma":[0.0000011444791,0.000009248008,0.00007648556,0.0000015136748,0.000003997872,0.0000054769644,0.0000035888575,0.998089,0.0010907,0.00060109736,0.00011539407,0.0000023442442],"about_ca_topic_score_codex":0.005379766,"about_ca_topic_score_gemma":0.005456376,"teacher_disagreement_score":0.005379766,"about_ca_system_score_codex":0.00040556453,"about_ca_system_score_gemma":0.00066558213,"threshold_uncertainty_score":0.010696888},"labels":[],"label_agreement":null},{"id":"W4402352450","doi":"10.1109/ijcnn60899.2024.10650050","title":"From Static to Dynamic: A Deeper, Faster, and Adaptive Language Modeling Approach","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"National Natural Science Foundation of China","keywords":"Computer science","score_opus":0.024774218982618966,"score_gpt":0.26469659444991483,"score_spread":0.23992237546729586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402352450","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03335012,0.0008364153,0.95402974,0.00070461887,0.00008181598,0.0000816308,0.00031746077,0.0075802268,0.0030180812],"genre_scores_gemma":[0.61237574,0.0008498139,0.3716016,0.0010083626,0.00014899891,0.0001923627,0.0015265013,0.0009806562,0.011315911],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99950516,0.00011223738,0.000026650554,0.0001827901,0.00009787818,0.000075336444],"domain_scores_gemma":[0.9993912,0.00020976251,0.000043079544,0.00017230767,0.00012529352,0.0000583043],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011287922,0.0015279497,0.0008926785,0.0010206802,0.00046515453,0.0011357827,0.0030754765,0.0011900323,0.0027785324],"category_scores_gemma":[0.0023784982,0.0008371231,0.0015391529,0.00093146873,0.0008555142,0.004740554,0.0019941328,0.0031417035,0.0015053957],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038002466,0.00020460624,0.0034766314,0.00016767991,0.00024058369,0.0001714319,0.00039214606,0.5228528,0.02167681,0.021399187,0.009627174,0.4194109],"study_design_scores_gemma":[0.0000149843945,0.000041487227,0.00019656112,0.000007157997,0.000035999896,0.00003694373,0.000018896562,0.9877101,0.0018444485,0.008312689,0.0017660363,0.000014577115],"about_ca_topic_score_codex":0.016461357,"about_ca_topic_score_gemma":0.030914636,"teacher_disagreement_score":0.016461357,"about_ca_system_score_codex":0.0014809311,"about_ca_system_score_gemma":0.002026899,"threshold_uncertainty_score":0.032731116},"labels":[],"label_agreement":null},{"id":"W4402402630","doi":"10.1109/ialp63756.2024.10661174","title":"The Power of Personalized Datasets: Advancing Chinese Composition Writing for Elementary School through Targeted Model Fine-Tuning","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Composition (language); Computer science; Power (physics); Mathematics education; Psychology; Linguistics; Physics","score_opus":0.016611748627520094,"score_gpt":0.30053838469815725,"score_spread":0.2839266360706372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402402630","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5084847,0.004574466,0.38431263,0.005818366,0.0015042938,0.0007488713,0.033091635,0.05101542,0.0104495585],"genre_scores_gemma":[0.7116577,0.00046075432,0.2218772,0.0007518083,0.0002089939,0.0006690076,0.060072657,0.0012314947,0.0030704401],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974631,0.0014236203,0.00012393208,0.0006761274,0.00020235653,0.00011080348],"domain_scores_gemma":[0.99334794,0.0036862406,0.00015802676,0.0018438332,0.00068138936,0.0002825238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005294409,0.0014787476,0.00077489205,0.0012794025,0.0009209974,0.0018629202,0.0022875594,0.0014680506,0.002248689],"category_scores_gemma":[0.020331528,0.0004925477,0.0013434547,0.0013011459,0.0005949813,0.0032700235,0.0022074948,0.0028019634,0.0017094518],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013144796,0.0016913427,0.049381893,0.0010097586,0.00084317935,0.0003969254,0.0018398543,0.3330589,0.0107916705,0.009594078,0.10116659,0.48891136],"study_design_scores_gemma":[0.00016414761,0.00012666237,0.0033421125,0.00006067852,0.00010121254,0.00004840247,0.00039795748,0.96285903,0.0049879197,0.010753224,0.017113157,0.00004548739],"about_ca_topic_score_codex":0.012963368,"about_ca_topic_score_gemma":0.025783405,"teacher_disagreement_score":0.012963368,"about_ca_system_score_codex":0.0013176297,"about_ca_system_score_gemma":0.0019274344,"threshold_uncertainty_score":0.027999878},"labels":[],"label_agreement":null},{"id":"W4402473640","doi":"10.1109/ccece59415.2024.10667245","title":"Adapting Large Language Models for Automatic Annotation of Radiology Reports for Metastases Detection","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Annotation; Artificial intelligence; Natural language processing","score_opus":0.03076148918633741,"score_gpt":0.2961688213673668,"score_spread":0.26540733218102935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402473640","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12274611,0.0025726813,0.8391684,0.0014346356,0.00041397626,0.0004558847,0.0072638504,0.023852123,0.0020923303],"genre_scores_gemma":[0.53651094,0.001259725,0.4317301,0.0006427054,0.0004124787,0.0007974526,0.023037335,0.0010958622,0.004513469],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99835396,0.00058244314,0.00016304804,0.00055785646,0.00022541032,0.000117357595],"domain_scores_gemma":[0.99140054,0.00615044,0.0005387645,0.00074406347,0.0010063123,0.00015980878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032215454,0.0016357765,0.00087193254,0.0028703858,0.00057110336,0.0017415283,0.0023181282,0.001365768,0.001215876],"category_scores_gemma":[0.009759118,0.00080801663,0.002280236,0.0017516863,0.00042217152,0.0024821304,0.0010989137,0.0022514015,0.0025809938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009515437,0.0008743239,0.026897239,0.00079266395,0.00066587527,0.00081533653,0.000991895,0.2902848,0.03128271,0.003382326,0.031664927,0.6113964],"study_design_scores_gemma":[0.00002747164,0.000057743855,0.0021091502,0.000023838611,0.00008376688,0.0001378218,0.00010117762,0.9865259,0.0044923555,0.002686092,0.0037206353,0.000034177992],"about_ca_topic_score_codex":0.01814032,"about_ca_topic_score_gemma":0.03272371,"teacher_disagreement_score":0.01814032,"about_ca_system_score_codex":0.0012799309,"about_ca_system_score_gemma":0.0017796083,"threshold_uncertainty_score":0.036069453},"labels":[],"label_agreement":null},{"id":"W4402475639","doi":"10.1109/ccece59415.2024.10667232","title":"Are Large Language Models General-Purpose Solvers for Dialogue Breakdown Detection? An Empirical Investigation","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Empirical research; Natural language processing; Epistemology; Philosophy","score_opus":0.04853789792381037,"score_gpt":0.30744230264120803,"score_spread":0.2589044047173977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402475639","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6614364,0.008733458,0.28803056,0.009009509,0.00065461773,0.0010541424,0.005004548,0.006607579,0.019469162],"genre_scores_gemma":[0.94170624,0.00068550505,0.050778594,0.0008258032,0.00016780906,0.0003480656,0.003718655,0.0004920583,0.0012772726],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.97610974,0.016511044,0.00079056656,0.0042762267,0.0016101063,0.0007023687],"domain_scores_gemma":[0.78619987,0.19154935,0.0042393734,0.011214137,0.0047513153,0.0020459276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029860236,0.0026666145,0.0017521115,0.0017804786,0.0010725153,0.0046666306,0.003168889,0.0025570844,0.007296931],"category_scores_gemma":[0.17801847,0.001018435,0.001450309,0.0017147657,0.001785818,0.010850475,0.0028476503,0.0060563423,0.0028521328],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00472718,0.0028574208,0.13500609,0.004132782,0.0014970051,0.00089650473,0.006765279,0.27178523,0.0047230795,0.02734004,0.044671822,0.49559754],"study_design_scores_gemma":[0.00021584763,0.0003502594,0.007713696,0.0002306978,0.00018793215,0.00031093738,0.0014595429,0.9595755,0.001098491,0.023554562,0.0052396385,0.00006290576],"about_ca_topic_score_codex":0.006361765,"about_ca_topic_score_gemma":0.007843548,"teacher_disagreement_score":0.029860236,"about_ca_system_score_codex":0.00194166,"about_ca_system_score_gemma":0.0022771757,"threshold_uncertainty_score":0.15791798},"labels":[],"label_agreement":null},{"id":"W4402527095","doi":"10.1007/978-3-031-71736-9_3","title":"Knowledge Acquisition Passage Retrieval: Corpus, Ranking Models, and Evaluation Resources","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Universiteit van Amsterdam; Canadian Institute of Steel Construction","keywords":"Computer science; Ranking (information retrieval); Information retrieval; Knowledge acquisition; Artificial intelligence; Natural language processing","score_opus":0.04013300437566199,"score_gpt":0.2820423754897383,"score_spread":0.24190937111407634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402527095","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06992168,0.04152049,0.78815824,0.004974718,0.0010372432,0.0010182617,0.013027574,0.014669875,0.06567193],"genre_scores_gemma":[0.34784713,0.014387498,0.54524606,0.0010959441,0.0019173549,0.0018628801,0.042580713,0.0028311752,0.042231146],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959246,0.0027645566,0.00020669417,0.00031761837,0.00068919675,0.00009738991],"domain_scores_gemma":[0.99076426,0.006298845,0.00023863312,0.00102482,0.001460413,0.00021300814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067872154,0.0009630036,0.001477993,0.0051076463,0.00080044696,0.0046475623,0.002014695,0.0012817517,0.011461036],"category_scores_gemma":[0.021552583,0.000626716,0.0007943882,0.007191401,0.000731349,0.005910929,0.0011676642,0.0014081629,0.009548597],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007373217,0.0005472432,0.002745532,0.0013456604,0.00023313107,0.00013914666,0.00030953702,0.0173808,0.0066759232,0.019910553,0.14251564,0.80745953],"study_design_scores_gemma":[0.00031666987,0.0011381144,0.014940784,0.0011149108,0.00092459854,0.0009800681,0.0005803603,0.7038818,0.028915271,0.070257924,0.1766045,0.00034505536],"about_ca_topic_score_codex":0.008024684,"about_ca_topic_score_gemma":0.00944288,"teacher_disagreement_score":0.011461036,"about_ca_system_score_codex":0.001664337,"about_ca_system_score_gemma":0.0022837122,"threshold_uncertainty_score":0.038341045},"labels":[],"label_agreement":null},{"id":"W4402552866","doi":"10.2196/59782","title":"Evaluating Medical Entity Recognition in Health Care: Entity Model Quantitative Study","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Computer science; Health care; Artificial intelligence; Data science; World Wide Web","score_opus":0.1318051687556064,"score_gpt":0.45194625194234095,"score_spread":0.3201410831867345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402552866","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.768256,0.009696441,0.18466112,0.0030982022,0.00028541012,0.00091188407,0.021856388,0.002111348,0.009123256],"genre_scores_gemma":[0.94492185,0.0007416057,0.039144333,0.00020054687,0.00009005626,0.00023308482,0.013872658,0.00007016659,0.0007257244],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9910602,0.004611746,0.00085731107,0.0016209545,0.0016262515,0.00022352632],"domain_scores_gemma":[0.9337635,0.054233935,0.0034518486,0.0031011566,0.0048512216,0.0005983273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015967084,0.0008298617,0.0006752766,0.0041901614,0.00049515493,0.00183384,0.0011676613,0.0012407971,0.002155337],"category_scores_gemma":[0.06513486,0.0001757853,0.0012175083,0.004719267,0.000673913,0.0039101834,0.0015367015,0.0012432714,0.0006969549],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014699973,0.000888341,0.47561315,0.0023004236,0.001699641,0.00032018346,0.0012807498,0.16642584,0.0024748882,0.008426621,0.019010205,0.32009],"study_design_scores_gemma":[0.00006403406,0.0008694557,0.13125144,0.000330289,0.0006198544,0.0006982972,0.0012088944,0.83521825,0.008416898,0.011260374,0.00994759,0.00011462377],"about_ca_topic_score_codex":0.0056549404,"about_ca_topic_score_gemma":0.004688838,"teacher_disagreement_score":0.015967084,"about_ca_system_score_codex":0.0023018804,"about_ca_system_score_gemma":0.0009847742,"threshold_uncertainty_score":0.08444303},"labels":[],"label_agreement":null},{"id":"W4402613910","doi":"10.23977/jaip.2024.070312","title":"Analysis of Errors at the Lexical Level in Post-editing for Medical Texts","year":2024,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Practice","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Natural language processing; Linguistics; Computer science; Philosophy","score_opus":0.1220433936815191,"score_gpt":0.40878029721595965,"score_spread":0.28673690353444053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402613910","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9772209,0.0007321806,0.009919331,0.00030627966,0.00029748812,0.0002639578,0.002009335,0.0008828007,0.00836772],"genre_scores_gemma":[0.970147,0.0003642437,0.019914106,0.000093965486,0.0000879947,0.00011166119,0.0029828444,0.0004526848,0.005845405],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9924493,0.0019048299,0.001374621,0.00087992096,0.0030818293,0.00030952532],"domain_scores_gemma":[0.85399276,0.10962173,0.010929494,0.0062318286,0.018262414,0.00096180063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032701867,0.00044671632,0.00036676278,0.004724605,0.0011186923,0.0024032462,0.0008159493,0.0009933074,0.0039325915],"category_scores_gemma":[0.07548419,0.0001936117,0.00023198857,0.0037052755,0.0010638706,0.0015635354,0.0014331916,0.00085430447,0.0024615629],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026346957,0.00072510744,0.20651397,0.0028857805,0.00017329954,0.0125262635,0.08537279,0.0017793153,0.059311334,0.0026715298,0.012658322,0.6127476],"study_design_scores_gemma":[0.00007844582,0.0016379483,0.63471717,0.0012599331,0.00039349057,0.0209806,0.05180219,0.024249617,0.14798732,0.00510107,0.11140128,0.0003909198],"about_ca_topic_score_codex":0.0019272089,"about_ca_topic_score_gemma":0.0024265076,"teacher_disagreement_score":0.004724605,"about_ca_system_score_codex":0.00055343815,"about_ca_system_score_gemma":0.00090954936,"threshold_uncertainty_score":0.017294645},"labels":[],"label_agreement":null},{"id":"W4402619319","doi":"10.1007/978-3-031-70242-6_32","title":"Intelligent Conversational Agent for Medical Information","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Pfizer (Canada)","funders":"","keywords":"Computer science; Intelligent agent; Human–computer interaction; Artificial intelligence; World Wide Web","score_opus":0.024133042930073012,"score_gpt":0.2646927494440479,"score_spread":0.2405597065139749,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402619319","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010542729,0.00698416,0.90419084,0.0025743344,0.0007787627,0.0002023533,0.00067821547,0.005545171,0.068503425],"genre_scores_gemma":[0.2410439,0.0041202363,0.6509607,0.00093558314,0.00058947713,0.00040907288,0.0020190922,0.00054331,0.099378705],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99971205,0.00012682141,0.000016303886,0.000041670457,0.000078573794,0.000024523373],"domain_scores_gemma":[0.9996563,0.00021288464,0.000015920186,0.000030184394,0.000050519953,0.000034121917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006900606,0.0004990372,0.0004465554,0.00047201882,0.00049663347,0.0012776192,0.00086090923,0.0010200236,0.010736145],"category_scores_gemma":[0.0014445927,0.00022511951,0.0004363691,0.0003491231,0.00022712564,0.0011875355,0.001072621,0.0009358801,0.004544776],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004216989,0.00035442604,0.00085532136,0.0007834445,0.00011736104,0.0008616312,0.0013213586,0.019973423,0.02814552,0.120306924,0.12700687,0.69985193],"study_design_scores_gemma":[0.00007388566,0.00021386807,0.0008262588,0.0002251681,0.00017266114,0.0015709953,0.00043980282,0.51035917,0.02025965,0.09442575,0.3713586,0.000074204494],"about_ca_topic_score_codex":0.0006446,"about_ca_topic_score_gemma":0.0009236222,"teacher_disagreement_score":0.010736145,"about_ca_system_score_codex":0.0003694769,"about_ca_system_score_gemma":0.0004929702,"threshold_uncertainty_score":0.03591597},"labels":[],"label_agreement":null},{"id":"W4402644087","doi":"10.1016/j.cose.2024.104120","title":"Entity and relation extractions for threat intelligence knowledge graphs","year":2024,"lang":"en","type":"article","venue":"Computers & Security","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Computer science; Knowledge graph; Relation (database); Natural language processing; Computer security; Knowledge management; Artificial intelligence; Data mining","score_opus":0.0321096837671894,"score_gpt":0.30012742349680904,"score_spread":0.2680177397296196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402644087","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016482692,0.0012624087,0.92157054,0.0014893688,0.00012639139,0.00082685763,0.029423708,0.021942094,0.006875852],"genre_scores_gemma":[0.14777295,0.0016787305,0.74961406,0.0006700404,0.00009954418,0.00051455054,0.09437517,0.000988651,0.004286247],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99796414,0.00027149316,0.00023508439,0.00073567004,0.0006736392,0.00012004531],"domain_scores_gemma":[0.9958281,0.0018802029,0.00048107057,0.0010461393,0.00062551646,0.00013903288],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013688952,0.0014887012,0.00065414247,0.0119389715,0.0011674684,0.002160147,0.0018758444,0.0014815899,0.0055067535],"category_scores_gemma":[0.012725236,0.0006018396,0.0024593521,0.0079914685,0.0008208581,0.0067982273,0.0033947018,0.0022434539,0.0027802864],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015638492,0.00031060693,0.01305882,0.0017539713,0.00034182242,0.0016765713,0.0012821857,0.06580542,0.0108432695,0.0882694,0.065468945,0.75103265],"study_design_scores_gemma":[0.00005833917,0.000096114694,0.00701095,0.00040546613,0.00036088226,0.001321116,0.0008801709,0.52743775,0.019872786,0.28045893,0.16199146,0.000106057705],"about_ca_topic_score_codex":0.0134359645,"about_ca_topic_score_gemma":0.031972136,"teacher_disagreement_score":0.0134359645,"about_ca_system_score_codex":0.0019382631,"about_ca_system_score_gemma":0.0027496896,"threshold_uncertainty_score":0.026715517},"labels":[],"label_agreement":null},{"id":"W4402646777","doi":"10.1007/978-3-031-70242-6_14","title":"Coherence Graphs: Bridging the Gap in Text Segmentation with Unsupervised Learning","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Bridging (networking); Computer science; Coherence (philosophical gambling strategy); Segmentation; Artificial intelligence; Natural language processing; Physics","score_opus":0.021355864718021943,"score_gpt":0.24323039815055045,"score_spread":0.2218745334325285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402646777","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009790624,0.0018937528,0.978119,0.00059309445,0.00014542161,0.0000916063,0.00092094205,0.0055016093,0.0029439384],"genre_scores_gemma":[0.17694849,0.002346912,0.79836994,0.0005325255,0.00071934256,0.00033291892,0.007976179,0.0042058113,0.0085679],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99789953,0.0007940719,0.00011841773,0.0006890807,0.0003804311,0.00011854032],"domain_scores_gemma":[0.9877995,0.009096676,0.0006528846,0.0013557711,0.00080493116,0.00029019086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021482888,0.0012777336,0.0014280119,0.0054716044,0.0012630522,0.002688181,0.002251782,0.0019211557,0.0074380157],"category_scores_gemma":[0.011048273,0.00093075755,0.001255039,0.0070442837,0.0016362414,0.0074739354,0.00334012,0.0021739346,0.004177483],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005229225,0.00018447089,0.0013594839,0.0007677986,0.00015718704,0.00015993394,0.0010778372,0.028381513,0.012399776,0.05429134,0.038537994,0.8621598],"study_design_scores_gemma":[0.000097424454,0.00014688645,0.0017972086,0.00016435838,0.00015725226,0.0001683045,0.0004696721,0.6199442,0.012929836,0.3244867,0.039563656,0.00007450547],"about_ca_topic_score_codex":0.0037127694,"about_ca_topic_score_gemma":0.0070299865,"teacher_disagreement_score":0.0074380157,"about_ca_system_score_codex":0.0008699359,"about_ca_system_score_gemma":0.0014906626,"threshold_uncertainty_score":0.024882615},"labels":[],"label_agreement":null},{"id":"W4402647202","doi":"10.1007/978-3-031-70242-6_37","title":"CoURAGE: A Framework to Evaluate RAG Systems","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"PricewaterhouseCoopers (Canada)","funders":"","keywords":"Courage; Computer science; Theology; Philosophy","score_opus":0.02741885497644407,"score_gpt":0.2844265836925258,"score_spread":0.25700772871608174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402647202","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009306331,0.00043367327,0.9636032,0.00034234882,0.00013618078,0.00017981094,0.00034617365,0.0034342534,0.022218004],"genre_scores_gemma":[0.38247687,0.0004553948,0.5914787,0.00018844697,0.00024063965,0.000681665,0.0007711065,0.0011604662,0.022546805],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99731123,0.0013395026,0.00010913748,0.000269986,0.00077324297,0.00019695847],"domain_scores_gemma":[0.99518675,0.0030565849,0.0004007358,0.00044536177,0.0006340059,0.00027654064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043092747,0.0014661351,0.00081729237,0.0020539712,0.0008229339,0.004077836,0.0023301255,0.0015693421,0.0151819885],"category_scores_gemma":[0.015766481,0.0003723875,0.00082799036,0.0011502873,0.0018889011,0.0050218883,0.0024177397,0.0017800167,0.0024238417],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024935443,0.00012085012,0.0015788991,0.00022476053,0.000099877645,0.00014229036,0.00028562083,0.21048847,0.00195211,0.6441998,0.01376514,0.12689285],"study_design_scores_gemma":[0.00003893808,0.00010020714,0.00031945854,0.000042206266,0.000038558388,0.000039163875,0.00009081401,0.75254023,0.0015476355,0.23254041,0.012669907,0.00003246496],"about_ca_topic_score_codex":0.0033674573,"about_ca_topic_score_gemma":0.0032234706,"teacher_disagreement_score":0.0151819885,"about_ca_system_score_codex":0.0013516118,"about_ca_system_score_gemma":0.0014761166,"threshold_uncertainty_score":0.05078882},"labels":[],"label_agreement":null},{"id":"W4402667062","doi":"10.18653/v1/2024.acl-long.20","title":"Confidence Under the Hood: An Investigation into the Confidence-Probability Alignment in Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Confidence interval; Artificial intelligence; Confidence distribution; Natural language processing; Statistics; Mathematics","score_opus":0.04172525238957015,"score_gpt":0.2864207796739621,"score_spread":0.24469552728439198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402667062","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5503473,0.0005932735,0.43784118,0.0017183597,0.0000704033,0.00030485913,0.00038031983,0.0012708569,0.0074734134],"genre_scores_gemma":[0.9611854,0.00009018019,0.037695024,0.00010929081,0.00003934349,0.00012981193,0.0003193737,0.00015900705,0.0002727341],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9586021,0.028720189,0.0021970014,0.004419213,0.0052452767,0.00081622007],"domain_scores_gemma":[0.4381969,0.4979623,0.025439337,0.02421485,0.011936827,0.0022497904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05227679,0.0011436163,0.00130887,0.0033049532,0.0018933095,0.0072391247,0.0019222522,0.0014459842,0.002782261],"category_scores_gemma":[0.38026574,0.0011841683,0.0013218177,0.003990561,0.0038668865,0.011077472,0.0064070746,0.0047061183,0.00049744674],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030116732,0.0009856493,0.43690425,0.0011210531,0.0011361482,0.0008466689,0.044573303,0.15058486,0.0078380685,0.13213328,0.0047904705,0.21607457],"study_design_scores_gemma":[0.00007994911,0.0006530913,0.05609411,0.00017301302,0.0001585691,0.0005646857,0.003814752,0.8229896,0.0044126487,0.10738974,0.0034618976,0.00020796976],"about_ca_topic_score_codex":0.0047425344,"about_ca_topic_score_gemma":0.003103728,"teacher_disagreement_score":0.05227679,"about_ca_system_score_codex":0.0025990286,"about_ca_system_score_gemma":0.0021030507,"threshold_uncertainty_score":0.2764694},"labels":[],"label_agreement":null},{"id":"W4402670397","doi":"10.18653/v1/2024.findings-acl.277","title":"SPIN: Sparsifying and Integrating Internal Neurons in Large Language Models for Text Classification","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Spin (aerodynamics); Artificial intelligence; Natural language processing; Physics","score_opus":0.05774731341226357,"score_gpt":0.31473679079301475,"score_spread":0.2569894773807512,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402670397","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020758016,0.0005699952,0.97180766,0.00036437117,0.00017278636,0.000069664224,0.00045835273,0.004918,0.000881152],"genre_scores_gemma":[0.37682799,0.00080679916,0.6075628,0.00053913816,0.00050513266,0.00036918587,0.0031864718,0.0010686864,0.0091338],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99957556,0.00013713093,0.000029723464,0.00011956313,0.000084821855,0.0000532394],"domain_scores_gemma":[0.9984363,0.00089762505,0.00008185856,0.0002763003,0.00021652278,0.00009135622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014057102,0.0012913629,0.0010174604,0.00080635137,0.00053869735,0.0010972923,0.0015372093,0.0014243869,0.003663696],"category_scores_gemma":[0.005001748,0.00061960384,0.0012002187,0.0012410956,0.0006450204,0.0022181228,0.0019703142,0.0030194977,0.0021767002],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059818366,0.0002226424,0.0017526649,0.0002066404,0.00025973146,0.00016851029,0.00029238762,0.1766212,0.014910673,0.010821871,0.020692404,0.7734531],"study_design_scores_gemma":[0.00002204301,0.000036595276,0.00016025739,0.00000910477,0.000028943563,0.000027120166,0.000028476496,0.9853232,0.002811319,0.010358517,0.001185318,0.000009189857],"about_ca_topic_score_codex":0.0062518204,"about_ca_topic_score_gemma":0.013335928,"teacher_disagreement_score":0.0062518204,"about_ca_system_score_codex":0.0005151724,"about_ca_system_score_gemma":0.0009583602,"threshold_uncertainty_score":0.012430847},"labels":[],"label_agreement":null},{"id":"W4402671301","doi":"10.18653/v1/2024.acl-long.785","title":"CausalGym: Benchmarking causal interpretability methods on linguistic tasks","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Cambridge; Institute for Catastrophic Loss Reduction","keywords":"Interpretability; Benchmarking; Computer science; Natural language processing; Artificial intelligence; Linguistics; Philosophy","score_opus":0.037292220696516734,"score_gpt":0.36924851090356064,"score_spread":0.3319562902070439,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402671301","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5905686,0.005395482,0.30901533,0.0037538945,0.0010752282,0.00196531,0.009256233,0.05858875,0.020381184],"genre_scores_gemma":[0.69610393,0.00066969765,0.28268945,0.0006315133,0.0001643463,0.0015568539,0.012388398,0.003576885,0.0022188933],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.985316,0.010242528,0.00083785306,0.0019554228,0.0012022426,0.0004460662],"domain_scores_gemma":[0.89241177,0.09326337,0.0014000488,0.00911885,0.0028206238,0.0009853515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020227287,0.002502538,0.0010252123,0.0032684603,0.0011867032,0.0026072704,0.0037351772,0.003771943,0.007169024],"category_scores_gemma":[0.09449512,0.0008075771,0.001757536,0.001986039,0.0019903271,0.0057522343,0.0042616813,0.0053985775,0.002017822],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040614204,0.0030192249,0.033927877,0.0039991597,0.0012855063,0.00046428572,0.0022553187,0.4736582,0.007763286,0.02774154,0.04239005,0.39943412],"study_design_scores_gemma":[0.0005142861,0.0004476337,0.0030369125,0.00009942176,0.000072596515,0.00007133448,0.0004017918,0.96668756,0.00463781,0.019427154,0.0045440663,0.00005946643],"about_ca_topic_score_codex":0.008425784,"about_ca_topic_score_gemma":0.011011472,"teacher_disagreement_score":0.020227287,"about_ca_system_score_codex":0.0024774133,"about_ca_system_score_gemma":0.003141847,"threshold_uncertainty_score":0.10697335},"labels":[],"label_agreement":null},{"id":"W4402671426","doi":"10.18653/v1/2024.arabicnlp-1.101","title":"WojoodNER 2024: The Second Arabic Named Entity Recognition Shared Task","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Social Sciences and Humanities Research Council of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Birzeit University","keywords":"Computer science; Arabic; Task (project management); Natural language processing; Artificial intelligence; Entity linking; Named-entity recognition; Speech recognition; Linguistics; Engineering; Knowledge base","score_opus":0.0318173123631445,"score_gpt":0.2448831006226526,"score_spread":0.2130657882595081,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402671426","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1436533,0.0052544335,0.32771775,0.0059660664,0.0069241147,0.003274189,0.2636211,0.16214877,0.08144023],"genre_scores_gemma":[0.17765068,0.000742588,0.2495154,0.0015085763,0.000830416,0.0025236437,0.5044807,0.007839337,0.054908626],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99698454,0.00089420634,0.00027599436,0.00089925795,0.00056511187,0.00038087618],"domain_scores_gemma":[0.99579144,0.0011936118,0.000107669366,0.0014153958,0.0009425332,0.0005492753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002961284,0.0037923886,0.002330086,0.0032478431,0.0028229915,0.0031174482,0.002992486,0.0038548068,0.03428072],"category_scores_gemma":[0.011239345,0.0009394796,0.0019005699,0.002574725,0.00060379884,0.006705962,0.00767647,0.0036038272,0.03404611],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019169489,0.0007218628,0.002644041,0.00087239675,0.00040258106,0.0012313249,0.000572054,0.004784942,0.016124496,0.004753967,0.5914008,0.37457457],"study_design_scores_gemma":[0.0013904236,0.0009844326,0.012718952,0.00035780735,0.00055164134,0.002378097,0.0020494338,0.21921082,0.08061546,0.033925634,0.64526916,0.00054808496],"about_ca_topic_score_codex":0.014859642,"about_ca_topic_score_gemma":0.019529052,"teacher_disagreement_score":0.03428072,"about_ca_system_score_codex":0.0011283358,"about_ca_system_score_gemma":0.0037900798,"threshold_uncertainty_score":0.11468041},"labels":[],"label_agreement":null},{"id":"W4402683840","doi":"10.18653/v1/2024.findings-acl.19","title":"Are self-explanations from Large Language Models faithful?","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Natural language processing; Language model; Artificial intelligence; Linguistics; Philosophy","score_opus":0.021926898320238954,"score_gpt":0.2564935360003454,"score_spread":0.23456663768010647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402683840","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21647282,0.00087435567,0.76800317,0.0046641505,0.0001606559,0.00018116024,0.001046524,0.0035793218,0.0050178426],"genre_scores_gemma":[0.9402697,0.00019532566,0.05596534,0.00077533483,0.00012471953,0.000096941214,0.0012167754,0.00043563204,0.00092018064],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9915558,0.0045280997,0.00049493357,0.0015567763,0.0014571783,0.00040714577],"domain_scores_gemma":[0.8850297,0.07908924,0.00898893,0.021858094,0.0038795199,0.0011545912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0116494745,0.0009817661,0.00088960223,0.0017268857,0.0008189693,0.0043297485,0.0019938014,0.0017066882,0.0035882555],"category_scores_gemma":[0.1111797,0.0010347387,0.001596327,0.0007458307,0.0026209522,0.0094000595,0.0031631526,0.004130177,0.00083497766],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024230254,0.0005465733,0.11141167,0.0014116907,0.0021872567,0.0013865108,0.008805436,0.18851863,0.015422087,0.28125298,0.014123061,0.37251106],"study_design_scores_gemma":[0.000098217286,0.00011199492,0.0063127805,0.00015490137,0.00016180881,0.0002715576,0.0004988066,0.51290184,0.005872414,0.4697159,0.0038056306,0.00009420346],"about_ca_topic_score_codex":0.003541689,"about_ca_topic_score_gemma":0.0044984673,"teacher_disagreement_score":0.0116494745,"about_ca_system_score_codex":0.0015438888,"about_ca_system_score_gemma":0.0017876227,"threshold_uncertainty_score":0.06160903},"labels":[],"label_agreement":null},{"id":"W4402684133","doi":"10.18653/v1/2024.findings-acl.85","title":"RIFF: Learning to Rephrase Inputs for Few-shot Fine-tuning of Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Alliance de recherche numérique du Canada; Alberta Machine Intelligence Institute; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Shot (pellet); Computer science; Language model; Artificial intelligence; One shot; Natural language processing; Programming language; Computer graphics (images); Mechanical engineering; Engineering; Materials science","score_opus":0.044489482314249255,"score_gpt":0.3024851167530286,"score_spread":0.2579956344387793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402684133","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055519648,0.001908609,0.89942783,0.00037070145,0.00029222635,0.0004026133,0.0015068145,0.03771743,0.002854168],"genre_scores_gemma":[0.3899243,0.0006173547,0.59073466,0.00080403703,0.00022688847,0.0009991413,0.0067398124,0.0023098856,0.007644024],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989723,0.00029529742,0.000047442867,0.0004812585,0.000108599874,0.000095207564],"domain_scores_gemma":[0.9969248,0.0018580661,0.00018009327,0.00056401367,0.0003339536,0.00013906938],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020222447,0.0026854554,0.0012431743,0.0010763904,0.0006325467,0.0011570548,0.0033669854,0.0022983733,0.00475127],"category_scores_gemma":[0.010248548,0.00075372116,0.0012011298,0.000677403,0.0006225351,0.0030669284,0.0015959259,0.0037018496,0.00520006],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052464625,0.0005277321,0.001963137,0.0007432772,0.00022383613,0.00037017767,0.0004822273,0.11352748,0.04718764,0.002601148,0.02217761,0.80967104],"study_design_scores_gemma":[0.00006673726,0.0003019041,0.00068241975,0.00004886711,0.00005184295,0.00021596966,0.00011225729,0.9654518,0.02260885,0.006250121,0.0041597076,0.00004947144],"about_ca_topic_score_codex":0.0034859579,"about_ca_topic_score_gemma":0.0073877876,"teacher_disagreement_score":0.00475127,"about_ca_system_score_codex":0.00081703585,"about_ca_system_score_gemma":0.0009816347,"threshold_uncertainty_score":0.015894592},"labels":[],"label_agreement":null},{"id":"W4402715036","doi":"10.18653/v1/2024.sigdial-1.26","title":"Coherence-based Dialogue Discourse Structure Extraction using Open-Source Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Mitacs","keywords":"Computer science; Coherence (philosophical gambling strategy); Natural language processing; Artificial intelligence; Extraction (chemistry); Linguistics; Physics; Philosophy; Chemistry","score_opus":0.041014228985578646,"score_gpt":0.3263863849262883,"score_spread":0.28537215594070964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402715036","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030023627,0.00088582025,0.93612605,0.0005346992,0.00014324834,0.0003081493,0.004419298,0.024645342,0.0029137314],"genre_scores_gemma":[0.26673457,0.00040969535,0.70925385,0.00015930625,0.0001551937,0.0006920813,0.017229209,0.001726996,0.003639101],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978496,0.00092331367,0.00013658326,0.00068608037,0.00030209668,0.00010233137],"domain_scores_gemma":[0.9962554,0.0022296605,0.0002932203,0.00043354827,0.00066650164,0.000121624966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019970995,0.0018580342,0.0007567526,0.0036392824,0.0011137659,0.0018851237,0.0012893466,0.0010885197,0.003489337],"category_scores_gemma":[0.008559921,0.00081917574,0.0012800973,0.0016340743,0.00062205765,0.0034998497,0.0026346608,0.0022820018,0.0031035508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004887262,0.00040325476,0.006087775,0.0014561104,0.00030539432,0.0005334707,0.004464996,0.0463951,0.05599121,0.013259738,0.03221998,0.8383942],"study_design_scores_gemma":[0.00013397807,0.00019016699,0.005579719,0.00018728158,0.00015348288,0.00025432502,0.0020866813,0.87518007,0.0339221,0.029497758,0.052690774,0.00012366261],"about_ca_topic_score_codex":0.0038141967,"about_ca_topic_score_gemma":0.008859846,"teacher_disagreement_score":0.0038141967,"about_ca_system_score_codex":0.00077778246,"about_ca_system_score_gemma":0.0020413257,"threshold_uncertainty_score":0.011673033},"labels":[],"label_agreement":null},{"id":"W4402716221","doi":"10.1109/cvpr52733.2024.00661","title":"CLiC: Concept Learning in Context","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Context (archaeology); Geology","score_opus":0.02669239094303115,"score_gpt":0.2753163558285381,"score_spread":0.24862396488550692,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402716221","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030313467,0.0032605682,0.96738565,0.0013483551,0.00041187642,0.00044542868,0.0020390858,0.0125128655,0.009564799],"genre_scores_gemma":[0.08973732,0.003011749,0.889296,0.0012543133,0.0005664913,0.0012890405,0.0061085154,0.0010789752,0.0076576495],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9975981,0.00081650296,0.00013184415,0.0007429386,0.00059653743,0.00011409115],"domain_scores_gemma":[0.99557245,0.0025952288,0.0001786636,0.0009165683,0.00045154864,0.00028558285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029443172,0.0019265376,0.0010528713,0.0036922144,0.0013631287,0.003937254,0.004481108,0.002223005,0.021820571],"category_scores_gemma":[0.013880602,0.00078406313,0.001733613,0.0031642716,0.00139104,0.010520044,0.006137717,0.0040802206,0.0065124403],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031750195,0.00023172563,0.0012300678,0.0015270391,0.00013012,0.00029292473,0.00052571757,0.022391012,0.0024614753,0.1359238,0.116362564,0.7186061],"study_design_scores_gemma":[0.00013593533,0.00013652218,0.0006584644,0.00048496772,0.00009947215,0.00059738866,0.00033687823,0.3359461,0.0058563026,0.39813405,0.25753865,0.00007518189],"about_ca_topic_score_codex":0.0049559907,"about_ca_topic_score_gemma":0.007498349,"teacher_disagreement_score":0.021820571,"about_ca_system_score_codex":0.0019635977,"about_ca_system_score_gemma":0.002824636,"threshold_uncertainty_score":0.07299709},"labels":[],"label_agreement":null},{"id":"W4402727495","doi":"10.1109/cvpr52733.2024.01432","title":"LMDrive: Closed-Loop End-to-End Driving with Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":121,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"End-to-end principle; Closed loop; Computer science; Loop (graph theory); End-user development; End user; Control theory (sociology); Engineering; Control engineering; Mathematics; Artificial intelligence; Operating system; Control (management)","score_opus":0.0139909095307899,"score_gpt":0.25804953909438444,"score_spread":0.24405862956359453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402727495","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029221853,0.00036976943,0.81402993,0.0002902857,0.00015978869,0.00039736106,0.001924189,0.14973627,0.00387046],"genre_scores_gemma":[0.39206308,0.00029311632,0.5812185,0.00062762364,0.000063168925,0.001089318,0.010241905,0.0057080993,0.008695171],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993061,0.00016057528,0.000038227983,0.00026216733,0.00015913176,0.00007374591],"domain_scores_gemma":[0.9988865,0.00060296594,0.000052356794,0.0001882636,0.00018050894,0.000089351415],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010753323,0.0016396185,0.00083679846,0.0005629132,0.000516558,0.0013353671,0.0033769882,0.0012102235,0.0057994216],"category_scores_gemma":[0.004807062,0.00066691363,0.0010970143,0.0002914156,0.00055956905,0.0026347472,0.0031416768,0.0018984736,0.0034295795],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015438594,0.001211011,0.00404621,0.0008070351,0.0004263383,0.0009453948,0.0012538393,0.29148462,0.03214652,0.012262896,0.0891996,0.5646727],"study_design_scores_gemma":[0.00007294382,0.00008564509,0.0001789398,0.000012722265,0.000015407835,0.000060417442,0.00006522971,0.978949,0.006040944,0.0056043672,0.008886194,0.000028238797],"about_ca_topic_score_codex":0.010071557,"about_ca_topic_score_gemma":0.01790442,"teacher_disagreement_score":0.010071557,"about_ca_system_score_codex":0.0006256117,"about_ca_system_score_gemma":0.0013256214,"threshold_uncertainty_score":0.02002585},"labels":[],"label_agreement":null},{"id":"W4402742293","doi":"10.1109/comst.2024.3465447","title":"Large Language Model (LLM) for Telecommunications: A Comprehensive Survey on Principles, Key Techniques, and Opportunities","year":2024,"lang":"en","type":"article","venue":"IEEE Communications Surveys & Tutorials","topic":"Topic Modeling","field":"Computer Science","cited_by":172,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; Western University; McGill University","funders":"","keywords":"Key (lock); Computer science; Telecommunications; Computer security","score_opus":0.2304285349274346,"score_gpt":0.3701655304340794,"score_spread":0.1397369955066448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402742293","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016398961,0.03385774,0.9517495,0.0033683076,0.00040465203,0.00014524738,0.000948277,0.0021720422,0.0057143993],"genre_scores_gemma":[0.13016436,0.11621287,0.7282618,0.0031741206,0.0030965116,0.0017543853,0.005521008,0.00195515,0.009859772],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99754286,0.0010317824,0.00023310311,0.00043358767,0.0006549985,0.000103698054],"domain_scores_gemma":[0.99520564,0.0035244473,0.00022923356,0.00050342537,0.00045476938,0.000082572886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034691868,0.0020315845,0.0018576187,0.0020964423,0.00063974573,0.004044163,0.0029565163,0.002081028,0.005377894],"category_scores_gemma":[0.009339301,0.0010956472,0.0024893272,0.002726416,0.0014193017,0.005517641,0.002678078,0.0047716475,0.0032541489],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012717294,0.0001529362,0.0015899333,0.0022586759,0.00033803852,0.000367132,0.00040431437,0.13726923,0.0023115673,0.27644134,0.030879667,0.54785997],"study_design_scores_gemma":[0.000028877088,0.00008899826,0.0004010573,0.000544879,0.00009324243,0.00026186323,0.00010205862,0.6639373,0.0012364786,0.24658616,0.08662462,0.00009444822],"about_ca_topic_score_codex":0.007847992,"about_ca_topic_score_gemma":0.006195855,"teacher_disagreement_score":0.007847992,"about_ca_system_score_codex":0.0021369623,"about_ca_system_score_gemma":0.0029990305,"threshold_uncertainty_score":0.018347025},"labels":[],"label_agreement":null},{"id":"W4402821735","doi":"10.1101/2024.09.23.614603","title":"Protein Language Models: Is Scaling Necessary?","year":2024,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; Polytechnique Montréal; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Scaling; Computer science; Mathematics; Geometry","score_opus":0.02002412734202891,"score_gpt":0.23131075610278704,"score_spread":0.21128662876075813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402821735","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17081815,0.021595953,0.6548549,0.074945286,0.003644729,0.0006271387,0.0057472754,0.022309719,0.045456897],"genre_scores_gemma":[0.6608684,0.009095579,0.29904804,0.007882149,0.0025985802,0.00081397255,0.007476837,0.0052362164,0.0069802706],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961869,0.0014413977,0.00025617995,0.0011336141,0.0007616995,0.00022027972],"domain_scores_gemma":[0.97381514,0.013389446,0.00085885084,0.008860558,0.0023689955,0.00070710975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010083112,0.0016420122,0.001383715,0.00091104244,0.00083667604,0.0029816183,0.0028443295,0.0018193265,0.0078043966],"category_scores_gemma":[0.07018971,0.0011175973,0.0010262238,0.0014715386,0.0015590184,0.0137327975,0.0034611234,0.0049728183,0.0055033388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085465756,0.0006352757,0.025963934,0.00089154835,0.0005034556,0.00040569506,0.0010982165,0.10687757,0.008734591,0.08948755,0.110311255,0.6542363],"study_design_scores_gemma":[0.00023981094,0.00024031024,0.0033814216,0.00037134765,0.00020151974,0.000352432,0.00056363636,0.5666887,0.0050062374,0.36543578,0.057423774,0.00009505339],"about_ca_topic_score_codex":0.0052671814,"about_ca_topic_score_gemma":0.005340844,"teacher_disagreement_score":0.010083112,"about_ca_system_score_codex":0.0010391289,"about_ca_system_score_gemma":0.0017744739,"threshold_uncertainty_score":0.053325236},"labels":[],"label_agreement":null},{"id":"W4402943559","doi":"10.1007/s11760-024-03560-z","title":"Integrating YOLO and WordNet for automated image object summarization","year":2024,"lang":"en","type":"article","venue":"Signal Image and Video Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Automatic summarization; WordNet; Computer science; Artificial intelligence; Computer vision; Object (grammar); Image (mathematics); Information retrieval; Natural language processing","score_opus":0.014880209860290998,"score_gpt":0.2875325919453473,"score_spread":0.2726523820850563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402943559","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046531975,0.004914659,0.84805775,0.001312635,0.0010657195,0.00086176634,0.014800299,0.07294469,0.009510575],"genre_scores_gemma":[0.2127614,0.002414143,0.7154241,0.0007338989,0.00087972987,0.0007316494,0.051620487,0.0030210994,0.012413551],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991148,0.00015653447,0.00011366857,0.0003106754,0.00017705178,0.0001272366],"domain_scores_gemma":[0.9985702,0.00041710795,0.00012698838,0.00018459158,0.0006016482,0.00009942406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013050243,0.002195114,0.0015021753,0.0073843473,0.0008649432,0.0025470438,0.0012199404,0.0013096861,0.0073776995],"category_scores_gemma":[0.0026894175,0.000562094,0.0012123326,0.0035472382,0.0003973659,0.003978453,0.0015999814,0.0012402514,0.008636748],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070150074,0.00039188055,0.0022956496,0.0012751123,0.00027934072,0.0002487073,0.00038956583,0.0067541916,0.051585443,0.0043999283,0.047979444,0.8836992],"study_design_scores_gemma":[0.00017759427,0.00070693105,0.0074925735,0.00038970544,0.0007087334,0.00046768508,0.0015543683,0.7702153,0.078027435,0.03129497,0.10878097,0.00018382593],"about_ca_topic_score_codex":0.007841233,"about_ca_topic_score_gemma":0.015576304,"teacher_disagreement_score":0.007841233,"about_ca_system_score_codex":0.0009483994,"about_ca_system_score_gemma":0.0017867424,"threshold_uncertainty_score":0.024680912},"labels":[],"label_agreement":null},{"id":"W4402977990","doi":"10.1145/3640310.3674091","title":"Text2VQL: Teaching a Model Query Language to Open-Source Language Models with ChatGPT","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; RDF query language; Query language; Open source; Programming language; Language model; Object Query Language; Natural language processing; Query by Example; Artificial intelligence; World Wide Web; Web search query; Information retrieval; Web query classification; Software; Search engine","score_opus":0.020833187557259894,"score_gpt":0.2853979916821693,"score_spread":0.2645648041249094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402977990","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048599085,0.0004061262,0.69330686,0.0022996636,0.00041824149,0.0002798306,0.01398118,0.27550003,0.008948144],"genre_scores_gemma":[0.14016971,0.0011875198,0.66451526,0.0027023982,0.00031780478,0.0012615328,0.07391356,0.09747185,0.018460397],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978841,0.0007664234,0.000181027,0.00048634995,0.0005657326,0.00011626415],"domain_scores_gemma":[0.99467045,0.00350632,0.00012445134,0.0009854075,0.0004988662,0.00021450665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033648915,0.0016155181,0.0008361656,0.0011683666,0.0006440342,0.0027069643,0.0037393884,0.0015310899,0.040832452],"category_scores_gemma":[0.0163659,0.0010453828,0.0016894195,0.0012909595,0.00079164054,0.007074228,0.005224042,0.003486125,0.018000744],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00061735115,0.00031936713,0.002078419,0.0014855867,0.00022301871,0.00060125336,0.0018137036,0.021211656,0.00948397,0.06993414,0.6523231,0.23990841],"study_design_scores_gemma":[0.00030846725,0.00013869687,0.00062713044,0.00029512966,0.000065537504,0.00040117902,0.00068700174,0.40829417,0.015870063,0.117075756,0.45611104,0.00012591417],"about_ca_topic_score_codex":0.005663076,"about_ca_topic_score_gemma":0.008774392,"teacher_disagreement_score":0.040832452,"about_ca_system_score_codex":0.0011569313,"about_ca_system_score_gemma":0.0016112116,"threshold_uncertainty_score":0.13659817},"labels":[],"label_agreement":null},{"id":"W4403090795","doi":"10.1101/2024.10.02.24314783","title":"Inconsistency detection in cancer data classification using explainable-AI","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital","funders":"","keywords":"Computer science; Artificial intelligence; Cancer detection; Cancer; Machine learning; Pattern recognition (psychology); Medicine; Internal medicine","score_opus":0.1801339199482145,"score_gpt":0.3615955237455872,"score_spread":0.18146160379737267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403090795","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0552337,0.0010861416,0.93911046,0.00057785085,0.00008782661,0.00011344438,0.0009853592,0.002302302,0.0005028865],"genre_scores_gemma":[0.5954217,0.00041182,0.3965655,0.00029917856,0.00026154213,0.00029217172,0.0055066706,0.00030457173,0.0009369587],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949414,0.0016581452,0.0004833077,0.0014807377,0.0011325937,0.00030378057],"domain_scores_gemma":[0.98221725,0.010458504,0.0025015941,0.0026186574,0.0019515442,0.00025246755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005738057,0.0012728136,0.0015015597,0.0059702476,0.0010382946,0.0024376393,0.0025689784,0.0017252185,0.001124497],"category_scores_gemma":[0.024960909,0.0004701344,0.0018140986,0.004181092,0.001031798,0.0024141714,0.0023833106,0.0022717568,0.00051361],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010537853,0.000396824,0.08734139,0.00082413666,0.0013175829,0.0010702352,0.0015058472,0.27923146,0.015098972,0.0163721,0.01110474,0.5846829],"study_design_scores_gemma":[0.000024444924,0.000057924764,0.0051931296,0.000033836615,0.000072579416,0.00015641491,0.00013079042,0.97123307,0.0032665222,0.01786735,0.0019347793,0.00002917405],"about_ca_topic_score_codex":0.0058386847,"about_ca_topic_score_gemma":0.0065835975,"teacher_disagreement_score":0.0059702476,"about_ca_system_score_codex":0.0012658873,"about_ca_system_score_gemma":0.0014531243,"threshold_uncertainty_score":0.030346155},"labels":[],"label_agreement":null},{"id":"W4403093099","doi":"10.1016/j.jml.2024.104573","title":"An embedded computational framework of memory: Accounting for the influence of semantic information in verbal short-term memory","year":2024,"lang":"en","type":"article","venue":"Journal of Memory and Language","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia; University of Manitoba; Université de Moncton","funders":"Natural Sciences and Engineering Research Council of Canada; Experimental Psychology Society","keywords":"Psychology; Short-term memory; Cognitive psychology; Semantic memory; Term (time); Verbal memory; Long-term memory; Cognition; Working memory; Neuroscience","score_opus":0.010621753145568912,"score_gpt":0.2855527774992688,"score_spread":0.27493102435369987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403093099","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11522187,0.00044366505,0.876363,0.00083282613,0.000076542914,0.00008518417,0.00013903806,0.00036620052,0.0064716786],"genre_scores_gemma":[0.85732603,0.00036726272,0.13923155,0.00017012747,0.00009212231,0.00019305512,0.0001842869,0.000078807956,0.0023567798],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99925107,0.00029265133,0.000049567465,0.00014475964,0.00018174201,0.00008011278],"domain_scores_gemma":[0.9966209,0.0018875878,0.00038284544,0.00070239994,0.00028839137,0.00011799112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015913162,0.00082068914,0.0008146918,0.000950648,0.00044449707,0.0028369008,0.0029302295,0.0011125639,0.003480876],"category_scores_gemma":[0.0075036106,0.00044900243,0.0020368116,0.000785175,0.0018630719,0.0066401316,0.0016480498,0.0016857554,0.00039437346],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036656097,0.00021943431,0.004636853,0.00030006387,0.00025352006,0.00043983312,0.0007795054,0.3303306,0.0109580485,0.54441357,0.0012900832,0.106012076],"study_design_scores_gemma":[0.000022846996,0.00012749158,0.0010055355,0.000022881231,0.00004400872,0.00014245899,0.000051673327,0.6545754,0.0017841525,0.3413595,0.00082906615,0.00003494379],"about_ca_topic_score_codex":0.0024889566,"about_ca_topic_score_gemma":0.0017621171,"teacher_disagreement_score":0.003480876,"about_ca_system_score_codex":0.001150609,"about_ca_system_score_gemma":0.0011926771,"threshold_uncertainty_score":0.011644721},"labels":[],"label_agreement":null},{"id":"W4403094786","doi":"10.22541/au.172797494.46569043/v1","title":"Evaluation of Local Explainability Methods in Turkish Text Classification Tasks","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Turkish; Computer science; Artificial intelligence; Natural language processing; Linguistics; Philosophy","score_opus":0.1800156534794118,"score_gpt":0.443886896661136,"score_spread":0.26387124318172417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403094786","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7698892,0.0046498803,0.21380582,0.0006507981,0.00017542581,0.00032672624,0.0013917092,0.005227704,0.0038827814],"genre_scores_gemma":[0.9374149,0.00039533232,0.056434497,0.00007586661,0.000094323106,0.00014359942,0.004045476,0.00022516615,0.001170753],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973654,0.0012097153,0.00019699772,0.0006940493,0.0003738488,0.00016004212],"domain_scores_gemma":[0.9844818,0.012262388,0.0007787074,0.0011268622,0.0009828489,0.00036743723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060807434,0.0017105101,0.0008505077,0.0026603497,0.0005899221,0.0011369176,0.0013161954,0.0016378867,0.002101815],"category_scores_gemma":[0.01724323,0.00023141043,0.0010982392,0.0015979771,0.0006609065,0.002591879,0.0015049938,0.0016892345,0.0005659784],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036513507,0.0010806137,0.04653588,0.0012207536,0.00096833816,0.00043928524,0.00087530306,0.26013806,0.012584439,0.00449491,0.007832164,0.6601788],"study_design_scores_gemma":[0.00014128695,0.00076371426,0.01371509,0.000057218404,0.00021228321,0.00015623849,0.0002905156,0.9712068,0.008270248,0.0039730347,0.0011722957,0.00004125892],"about_ca_topic_score_codex":0.0042770165,"about_ca_topic_score_gemma":0.006070731,"teacher_disagreement_score":0.0060807434,"about_ca_system_score_codex":0.0010855932,"about_ca_system_score_gemma":0.00077030604,"threshold_uncertainty_score":0.032158434},"labels":[],"label_agreement":null},{"id":"W4403170239","doi":"10.24251/hicss.2023.904","title":"Narrating Causal Graphs with Large Language Models","year":2024,"lang":"en","type":"article","venue":"Proceedings of the ... Annual Hawaii International Conference on System Sciences/Proceedings of the Annual Hawaii International Conference on System Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Lakehead University","funders":"","keywords":"Computer science; Linguistics; Natural language processing; Philosophy","score_opus":0.040031523653819165,"score_gpt":0.2926707159375463,"score_spread":0.2526391922837271,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403170239","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052632537,0.00096626027,0.92190886,0.0030167873,0.00022137722,0.0003239222,0.0059006712,0.009218282,0.0058112624],"genre_scores_gemma":[0.5610381,0.0007416069,0.4140881,0.0010238114,0.00021356354,0.0006635462,0.015683597,0.00089569564,0.005652075],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980736,0.0010945156,0.000095514246,0.00044643157,0.00023027352,0.00005962861],"domain_scores_gemma":[0.9754719,0.021871896,0.00055789994,0.0012390005,0.0006712313,0.00018802637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027168763,0.001133007,0.0004731164,0.00199206,0.0007377284,0.0016940573,0.0017846606,0.0014592487,0.006907445],"category_scores_gemma":[0.025699323,0.00064360554,0.0016645637,0.0013828876,0.0010944792,0.0049483636,0.002099894,0.0024785132,0.0016910615],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004014729,0.00032090305,0.006260407,0.001289547,0.00027606657,0.0016269409,0.0021927736,0.58895785,0.0046859225,0.108474575,0.03862507,0.24688844],"study_design_scores_gemma":[0.000052009316,0.000034001234,0.0003003214,0.000055375192,0.000034518835,0.00013609823,0.00019259944,0.88660717,0.0017187652,0.100874744,0.00997227,0.00002207116],"about_ca_topic_score_codex":0.0063693165,"about_ca_topic_score_gemma":0.012390417,"teacher_disagreement_score":0.006907445,"about_ca_system_score_codex":0.0014585019,"about_ca_system_score_gemma":0.0010194067,"threshold_uncertainty_score":0.023107648},"labels":[],"label_agreement":null},{"id":"W4403221620","doi":"10.1145/3640457.3688185","title":"EmbSum: Leveraging the Summarization Capabilities of Large Language Models for Content-Based Recommendations","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Content (measure theory); Natural language processing; Language model; Information retrieval","score_opus":0.06651139834127164,"score_gpt":0.29536027343680654,"score_spread":0.22884887509553492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403221620","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013551787,0.002266975,0.97188675,0.0003335035,0.00014912759,0.00013669275,0.0015506957,0.0090239495,0.0011006844],"genre_scores_gemma":[0.2989089,0.0022705146,0.67290264,0.00065959955,0.00055419194,0.00052681705,0.012390156,0.0011635892,0.010623628],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992405,0.00026386752,0.00006315354,0.00022154827,0.00015485927,0.000056053905],"domain_scores_gemma":[0.9979539,0.0012591848,0.0001062839,0.00030869778,0.00029328416,0.00007875724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014040346,0.0018228395,0.0012964513,0.0017737681,0.00045829325,0.0012363155,0.0017935863,0.0013587931,0.002905613],"category_scores_gemma":[0.005224166,0.00067944615,0.0012575717,0.0015832802,0.00033884714,0.0029758418,0.0012504697,0.0023820593,0.0041515906],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004426402,0.00035866987,0.003060664,0.00057095644,0.0004531103,0.00030503265,0.00039745457,0.20087695,0.016340453,0.0066432357,0.022202179,0.7483487],"study_design_scores_gemma":[0.000030274281,0.00009305279,0.00037086668,0.000020409701,0.0000705869,0.00007516432,0.000028282206,0.9865796,0.0033397793,0.0054738116,0.00389376,0.000024343612],"about_ca_topic_score_codex":0.010853782,"about_ca_topic_score_gemma":0.025995225,"teacher_disagreement_score":0.010853782,"about_ca_system_score_codex":0.00061223155,"about_ca_system_score_gemma":0.0010051721,"threshold_uncertainty_score":0.021581173},"labels":[],"label_agreement":null},{"id":"W4403246668","doi":"10.5617/dhnbpub.11184","title":"Name the Name","year":2020,"lang":"en","type":"article","venue":"Digital Humanities in the Nordic and Baltic Countries Publications","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Library of Parliament","funders":"","keywords":"Linguistics; Philosophy","score_opus":0.035421281463243935,"score_gpt":0.2351190811877139,"score_spread":0.19969779972446997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403246668","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020385631,0.0024203688,0.004185278,0.012831456,0.016524477,0.0002070256,0.008659771,0.0045139017,0.94861925],"genre_scores_gemma":[0.011440345,0.002334429,0.0018188927,0.008245537,0.0031036262,0.00011634992,0.005822111,0.0014771572,0.9656416],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99884063,0.00016524321,0.00009855536,0.00031371528,0.00037333678,0.0002085152],"domain_scores_gemma":[0.9970528,0.0003384634,0.00021935796,0.000593838,0.00091666175,0.00087879854],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001032141,0.00078142306,0.00085277914,0.0015074542,0.001835923,0.005396512,0.0012585699,0.0021481293,0.7197924],"category_scores_gemma":[0.0054539335,0.00029448705,0.00054839905,0.0017731502,0.0006451181,0.0059476965,0.003645519,0.0015460742,0.6920558],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000050833678,0.000024729814,0.0007204395,0.00020238713,0.000007863912,0.00019738349,0.00017745575,0.00003345213,0.0005576823,0.0057074917,0.8361483,0.15617195],"study_design_scores_gemma":[0.000003063803,0.000009842044,0.0002540944,0.000029840196,0.0000026647178,0.00018098381,0.00005741698,0.000019661116,0.000110798006,0.00059555174,0.9987312,0.0000049509536],"about_ca_topic_score_codex":0.001158334,"about_ca_topic_score_gemma":0.0016554593,"teacher_disagreement_score":0.7197924,"about_ca_system_score_codex":0.0008611337,"about_ca_system_score_gemma":0.0013942935,"threshold_uncertainty_score":0.39968204},"labels":[],"label_agreement":null},{"id":"W4403413296","doi":"10.1145/3674805.3686688","title":"Negative Results of Image Processing for Identifying Duplicate Questions on Stack Overflow","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Stack (abstract data type); Image processing; Image (mathematics); Artificial intelligence; Pattern recognition (psychology); Data mining; Information retrieval; Programming language","score_opus":0.05917570804844958,"score_gpt":0.3451050293138652,"score_spread":0.28592932126541565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403413296","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73585194,0.0034343444,0.22729272,0.003577664,0.0012059725,0.0008106537,0.003037025,0.011872734,0.012916954],"genre_scores_gemma":[0.8880278,0.00050181145,0.102426425,0.0009593415,0.00029953447,0.00019603806,0.0027388232,0.0008266177,0.0040235813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99066657,0.003677283,0.00078813295,0.0017912469,0.0025379215,0.0005388587],"domain_scores_gemma":[0.8900293,0.084300496,0.004397803,0.009365924,0.010640792,0.0012656513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01006278,0.0019687065,0.0012820477,0.0032482892,0.0013383471,0.0025554553,0.0014510019,0.0024886192,0.0028963154],"category_scores_gemma":[0.0942084,0.00041962773,0.0015662407,0.001597501,0.0017548307,0.0053683934,0.0020736824,0.002293613,0.001878521],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005555464,0.0012269715,0.12724662,0.0038841723,0.0007996926,0.003436627,0.008579237,0.040176775,0.08516488,0.0099951,0.045374405,0.66856],"study_design_scores_gemma":[0.00020023811,0.0016187401,0.07479171,0.0006921315,0.0008095712,0.0039864834,0.005382243,0.6286782,0.21049887,0.027547492,0.045341685,0.0004525524],"about_ca_topic_score_codex":0.0058020316,"about_ca_topic_score_gemma":0.0046005948,"teacher_disagreement_score":0.01006278,"about_ca_system_score_codex":0.0013049733,"about_ca_system_score_gemma":0.0014407337,"threshold_uncertainty_score":0.05321771},"labels":[],"label_agreement":null},{"id":"W4403413357","doi":"10.1145/3674805.3686695","title":"A Transformer-based Approach for Augmenting Software Engineering Chatbots Datasets","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of Calgary","funders":"","keywords":"Computer science; Transformer; Software; Software engineering; Programming language; Engineering; Electrical engineering","score_opus":0.024969716477097396,"score_gpt":0.2452655686757533,"score_spread":0.2202958521986559,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403413357","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20222917,0.001341841,0.67269206,0.0018422902,0.00047915132,0.0026356438,0.04638437,0.06622052,0.0061748936],"genre_scores_gemma":[0.31800336,0.00031881512,0.5411086,0.00040789758,0.00015498739,0.0021342444,0.1347031,0.0009139775,0.0022550153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99463105,0.0023827858,0.0005897032,0.0012107105,0.0009574996,0.00022835586],"domain_scores_gemma":[0.984743,0.0070655956,0.000732946,0.0041942606,0.0028416715,0.00042244786],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054605436,0.0013830853,0.0008441343,0.005770696,0.0010773963,0.0014542727,0.00263015,0.0011459603,0.0028983431],"category_scores_gemma":[0.02361804,0.0004655928,0.0017631248,0.0034611744,0.0007997101,0.0038148148,0.0044226865,0.0020009608,0.002394675],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014921479,0.0017998655,0.028010447,0.0021231193,0.00033340248,0.0012816595,0.0029951818,0.03826676,0.050684176,0.01172154,0.09010365,0.771188],"study_design_scores_gemma":[0.0003094513,0.0009809936,0.020118227,0.0002171578,0.00025217258,0.001354344,0.0026494851,0.7975005,0.039833944,0.025174962,0.111424424,0.0001844211],"about_ca_topic_score_codex":0.00439075,"about_ca_topic_score_gemma":0.009272046,"teacher_disagreement_score":0.005770696,"about_ca_system_score_codex":0.0011004499,"about_ca_system_score_gemma":0.0017046916,"threshold_uncertainty_score":0.02887845},"labels":[],"label_agreement":null},{"id":"W4403420631","doi":"10.1109/mipr62202.2024.00035","title":"FastLearn: A Rapid Learning Agent for Chat Models to Acquire Latest Knowledge","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Knowledge management; Human–computer interaction","score_opus":0.07826378248875293,"score_gpt":0.30031898245180627,"score_spread":0.22205519996305334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403420631","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034069445,0.0008553086,0.7869918,0.0007647119,0.00046080633,0.0009722717,0.002000331,0.16871262,0.0051727504],"genre_scores_gemma":[0.26425818,0.0004144996,0.7141084,0.0010759267,0.00018500422,0.0020777981,0.0047824737,0.0042375377,0.008860124],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983261,0.00066773157,0.00012275262,0.0005263787,0.0002477598,0.000109317014],"domain_scores_gemma":[0.99280256,0.0043706265,0.00034838828,0.0012564065,0.0008683628,0.0003536482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004687005,0.0024221588,0.0013091785,0.0010661504,0.000665234,0.0015950918,0.005018604,0.0022543308,0.011645868],"category_scores_gemma":[0.022946363,0.001026548,0.0014299987,0.0004942978,0.00065313827,0.006063229,0.0035321047,0.0043303217,0.005991051],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002034264,0.0015587116,0.009957989,0.0016028696,0.0005377042,0.00061302807,0.0019461705,0.13884136,0.019843716,0.010756228,0.07429424,0.7380138],"study_design_scores_gemma":[0.00017723454,0.0002204733,0.000274238,0.000043724776,0.00006775429,0.00007417003,0.000107499,0.97167856,0.007411424,0.005874594,0.0140174655,0.00005277314],"about_ca_topic_score_codex":0.0036853638,"about_ca_topic_score_gemma":0.008065664,"teacher_disagreement_score":0.011645868,"about_ca_system_score_codex":0.0010331233,"about_ca_system_score_gemma":0.0018502894,"threshold_uncertainty_score":0.038959384},"labels":[],"label_agreement":null},{"id":"W4403443673","doi":"10.48550/arxiv.2410.08821","title":"DeepNote: Note-Centric Deep Retrieval-Augmented Generation","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Cognitive psychology; Computer science; Labrador Retriever; Psychology; Information retrieval; Computer vision; Medicine","score_opus":0.07465491100708617,"score_gpt":0.1970579504967366,"score_spread":0.12240303948965044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403443673","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01745653,0.0011106163,0.949099,0.0005505678,0.0001752482,0.0002923129,0.0015666249,0.026231665,0.003517436],"genre_scores_gemma":[0.34271675,0.0005792251,0.63318497,0.0013820963,0.00019765027,0.00062608736,0.008997792,0.0023850498,0.009930375],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985298,0.00051346887,0.00007175696,0.0004930464,0.00026762384,0.0001242663],"domain_scores_gemma":[0.99633175,0.002023805,0.00012035073,0.0010514365,0.00033220806,0.00014050299],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023549283,0.0015147481,0.001135861,0.0010965399,0.0006297133,0.0017395498,0.004198267,0.0019873062,0.009830876],"category_scores_gemma":[0.008683017,0.0007755015,0.0013368661,0.0010079253,0.0011943949,0.0048500244,0.005465751,0.0031530873,0.0038565707],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006999103,0.0003961397,0.0030202768,0.0008019806,0.00020975515,0.00037688654,0.0009116389,0.11531952,0.026490139,0.022367535,0.05138737,0.7780188],"study_design_scores_gemma":[0.00016900143,0.00017398705,0.0004261224,0.000039994855,0.00006420232,0.00022397553,0.0001327435,0.9301239,0.0120218,0.04171049,0.01485873,0.00005506692],"about_ca_topic_score_codex":0.0037495808,"about_ca_topic_score_gemma":0.00971704,"teacher_disagreement_score":0.009830876,"about_ca_system_score_codex":0.00095698296,"about_ca_system_score_gemma":0.0016083461,"threshold_uncertainty_score":0.03288752},"labels":[],"label_agreement":null},{"id":"W4403582839","doi":"10.1145/3627673.3679534","title":"Aligning Query Representation with Rewritten Query and Relevance Judgments in Conversational Search","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Beijing Jiaotong University","keywords":"Computer science; Query expansion; Relevance (law); Sargable; Information retrieval; Query optimization; Query language; Web search query; Representation (politics); Relevance feedback; Web query classification; RDF query language; Query by Example; Search engine; Natural language processing; Artificial intelligence; Image retrieval","score_opus":0.03131823917913042,"score_gpt":0.28764364267091536,"score_spread":0.25632540349178495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403582839","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17453912,0.0034742123,0.806131,0.0007477577,0.00013912574,0.0005207356,0.0010794287,0.009861764,0.0035068528],"genre_scores_gemma":[0.7269532,0.00094729045,0.26071474,0.000558916,0.00021853615,0.00033973827,0.0045075733,0.00049305777,0.0052668923],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972626,0.0010131309,0.00020324875,0.0008257155,0.0005093938,0.00018588542],"domain_scores_gemma":[0.9969015,0.0016854613,0.00018233062,0.00061549613,0.00050447154,0.000110752604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025465714,0.001131371,0.0012405625,0.0015288921,0.0005763392,0.0012490847,0.0016595213,0.0012628735,0.0016142519],"category_scores_gemma":[0.010073887,0.00048693188,0.0009608516,0.00149531,0.0006836307,0.0037532505,0.0017664959,0.0016215368,0.001637493],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010581093,0.00094384025,0.006799415,0.0008305497,0.00029063397,0.00045019205,0.0016968988,0.11919914,0.06489132,0.0068152403,0.016375128,0.7806495],"study_design_scores_gemma":[0.00006818787,0.00038965262,0.0022328375,0.000023505929,0.00010719142,0.00041497673,0.00032911246,0.9680053,0.015642723,0.0062928805,0.006407066,0.0000866216],"about_ca_topic_score_codex":0.013958271,"about_ca_topic_score_gemma":0.013366507,"teacher_disagreement_score":0.013958271,"about_ca_system_score_codex":0.0009187549,"about_ca_system_score_gemma":0.0017431024,"threshold_uncertainty_score":0.027754068},"labels":[],"label_agreement":null},{"id":"W4403603005","doi":"10.2196/60164","title":"Health Care Language Models and Their Fine-Tuning for Information Extraction: Scoping Review","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Terminology; Computer science; Health care; Data extraction; Scopus; Unified Medical Language System; Knowledge management; MEDLINE; Data science; Artificial intelligence; Linguistics","score_opus":0.03758902317528386,"score_gpt":0.3692044030291242,"score_spread":0.3316153798538404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403603005","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011024368,0.9801501,0.009237923,0.0020704092,0.00040086586,0.003907403,0.0012504081,0.00013146485,0.0017490393],"genre_scores_gemma":[0.02023633,0.93441886,0.02911043,0.0013550697,0.00025176667,0.012714166,0.0015372036,0.00008496003,0.00029108068],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.90388316,0.046831563,0.033443365,0.004245263,0.010790006,0.0008065731],"domain_scores_gemma":[0.5391566,0.3928559,0.02897199,0.011696364,0.026536848,0.00078239926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.12800866,0.0031038313,0.009082157,0.036969233,0.0019835653,0.008463203,0.005672395,0.0035029126,0.007735293],"category_scores_gemma":[0.41426697,0.0024256087,0.014661478,0.02915619,0.0033273436,0.011128219,0.006854994,0.0032068929,0.0013612743],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014475876,0.000047102807,0.0009752969,0.78950334,0.0045266547,0.00009618213,0.0012299625,0.0007993666,0.00016511847,0.003497307,0.003698736,0.19531623],"study_design_scores_gemma":[0.00006248623,0.00006752701,0.0007194835,0.95956945,0.010749741,0.00010904476,0.0005783004,0.00042419962,0.0002690257,0.0022445328,0.02515793,0.000048418693],"about_ca_topic_score_codex":0.012673198,"about_ca_topic_score_gemma":0.017525831,"teacher_disagreement_score":0.12800866,"about_ca_system_score_codex":0.010761924,"about_ca_system_score_gemma":0.042888146,"threshold_uncertainty_score":0.67698264},"labels":[],"label_agreement":null},{"id":"W4403606135","doi":"10.3390/info15100634","title":"Promptology: Enhancing Human–AI Interaction in Large Language Models","year":2024,"lang":"en","type":"article","venue":"Information","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Natural language processing; Linguistics; Cognitive science; Psychology; Artificial intelligence; Philosophy","score_opus":0.015260533042689867,"score_gpt":0.29598498525756756,"score_spread":0.2807244522148777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403606135","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043650594,0.00039621358,0.9347134,0.0013687338,0.00012665571,0.0005553447,0.00025268033,0.01046688,0.008469389],"genre_scores_gemma":[0.31262827,0.0003811892,0.68016285,0.00043266974,0.00006737391,0.0005725963,0.000654429,0.0010192994,0.0040813033],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9914978,0.0062735644,0.0002896651,0.00078742893,0.000962835,0.00018861885],"domain_scores_gemma":[0.95997655,0.03144119,0.0011995282,0.0048441244,0.001672399,0.00086620176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009528375,0.0011917265,0.0006704777,0.0015103179,0.0010840069,0.004205767,0.0019181472,0.0017734641,0.008501571],"category_scores_gemma":[0.054497946,0.0005611318,0.0009304579,0.00083311246,0.0024826252,0.011110864,0.008288137,0.0022446571,0.0021121455],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010432849,0.0009367401,0.008750371,0.002644231,0.0001317169,0.0010631209,0.06809375,0.036942545,0.04191613,0.13380961,0.020324526,0.684344],"study_design_scores_gemma":[0.00039549993,0.0021170885,0.0046894965,0.0009862023,0.0002786399,0.0018896498,0.021282163,0.35918847,0.048522413,0.2830295,0.27722275,0.00039809998],"about_ca_topic_score_codex":0.00093873515,"about_ca_topic_score_gemma":0.001983118,"teacher_disagreement_score":0.009528375,"about_ca_system_score_codex":0.0010952136,"about_ca_system_score_gemma":0.0023689193,"threshold_uncertainty_score":0.050391436},"labels":[],"label_agreement":null},{"id":"W4403659146","doi":"10.1007/s11227-024-06499-7","title":"Enhancing Chinese comprehension and reasoning for large language models: an efficient LoRA fine-tuning and tree of thoughts framework","year":2024,"lang":"en","type":"article","venue":"The Journal of Supercomputing","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Sichuan Province Science and Technology Support Program","keywords":"Computer science; Comprehension; Tree (set theory); Artificial intelligence; Natural language processing; Theoretical computer science; Programming language","score_opus":0.017503935983843675,"score_gpt":0.2834888811719226,"score_spread":0.2659849451880789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403659146","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03251089,0.00024609515,0.94708645,0.00044695375,0.000073164774,0.0001816434,0.00048490838,0.015220597,0.0037493887],"genre_scores_gemma":[0.34943455,0.00015171211,0.6434933,0.00023513309,0.000088681205,0.00021805974,0.0013138988,0.0010054251,0.004059264],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987453,0.0003986592,0.000109307446,0.00032359836,0.0002793282,0.0001437521],"domain_scores_gemma":[0.9967218,0.0014474767,0.00011371352,0.0008047121,0.0007127628,0.00019947448],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015769731,0.00089558196,0.0013459054,0.0014900224,0.0010397716,0.0020362495,0.0024456193,0.0008341133,0.008616065],"category_scores_gemma":[0.006508486,0.00054441614,0.0020324383,0.0011268525,0.000802683,0.0046773762,0.002882354,0.0021850294,0.0024504412],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006464237,0.00061548565,0.0031002488,0.0004610073,0.00021141524,0.00049842737,0.001597068,0.06002975,0.027346227,0.0673518,0.023958512,0.81418365],"study_design_scores_gemma":[0.00005366978,0.00006642627,0.0002931091,0.000015346486,0.00009613193,0.00007780637,0.00021201825,0.9327101,0.0089239655,0.052014172,0.0055050203,0.0000322307],"about_ca_topic_score_codex":0.008757343,"about_ca_topic_score_gemma":0.014413259,"teacher_disagreement_score":0.008757343,"about_ca_system_score_codex":0.0012114841,"about_ca_system_score_gemma":0.0032128524,"threshold_uncertainty_score":0.028823614},"labels":[],"label_agreement":null},{"id":"W4403666342","doi":"10.48550/arxiv.2409.08212","title":"Adaptive Language-Guided Abstraction from Contrastive Explanations","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Defense Science and Engineering Graduate; Open Philanthropy Project; National Science Foundation","keywords":"Abstraction; Computer science; Programming language; Linguistics; Natural language processing; Contrastive analysis; Artificial intelligence; Philosophy; Epistemology","score_opus":0.10488071759278206,"score_gpt":0.2123155390683983,"score_spread":0.10743482147561625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403666342","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013361133,0.000102551945,0.98430264,0.00024696134,0.000017747705,0.000033614953,0.000075117234,0.00095380854,0.0009064218],"genre_scores_gemma":[0.5927497,0.0002378075,0.4029053,0.0002684826,0.00004475888,0.0002351338,0.00038914065,0.00028052146,0.0028891903],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99933594,0.0002607787,0.000029882764,0.00015764305,0.00015525958,0.00006051066],"domain_scores_gemma":[0.99777454,0.0014373226,0.00019627575,0.00040116272,0.00011556638,0.000075064985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000981781,0.0010520231,0.00053103286,0.00056939997,0.00033205267,0.0006534502,0.001943198,0.00091090583,0.0024736207],"category_scores_gemma":[0.0060031144,0.00048904814,0.0012493344,0.0003425737,0.0013553036,0.0022110718,0.0021640768,0.0021600872,0.00041373438],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034835056,0.00019448268,0.002757328,0.0004659534,0.00018637834,0.0007073959,0.001607527,0.47247764,0.02594993,0.24069521,0.006404288,0.24820553],"study_design_scores_gemma":[0.00003060333,0.00006289206,0.00025077513,0.000019657367,0.000024351355,0.00009924675,0.000050139388,0.8627874,0.004862154,0.12875235,0.0030413428,0.000019003199],"about_ca_topic_score_codex":0.00173622,"about_ca_topic_score_gemma":0.003344126,"teacher_disagreement_score":0.0024736207,"about_ca_system_score_codex":0.0009137601,"about_ca_system_score_gemma":0.0008846207,"threshold_uncertainty_score":0.008275092},"labels":[],"label_agreement":null},{"id":"W4403745910","doi":"10.18280/isi.290538","title":"Text Summarization: A Bibliometric Study and Systematic Literature Review","year":2024,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Bibliometrics; Information retrieval; Systematic review; Computer science; Data science; Library science; MEDLINE; Political science","score_opus":0.017356563989845582,"score_gpt":0.2589526443469761,"score_spread":0.24159608035713054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403745910","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16938224,0.77277094,0.013307091,0.0055881054,0.0008210647,0.01209938,0.01786148,0.00028686612,0.007882837],"genre_scores_gemma":[0.560276,0.3700401,0.03990614,0.0016645209,0.0011713094,0.013725295,0.0117152445,0.00014578048,0.0013556227],"study_design_codex":"systematic_review","study_design_gemma":"not_applicable","domain_scores_codex":[0.88832927,0.043830905,0.037622333,0.005785114,0.023213265,0.0012190461],"domain_scores_gemma":[0.6622913,0.2513903,0.040802266,0.0075402283,0.036000594,0.0019753545],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.07080765,0.0015280851,0.00772173,0.13822484,0.0026956021,0.0077532586,0.002429595,0.002424671,0.00469193],"category_scores_gemma":[0.24829233,0.0011959443,0.0069041033,0.14225988,0.0021441358,0.009590877,0.0043797637,0.0011620866,0.00066769554],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015969595,0.0005282944,0.06736104,0.49477264,0.022179687,0.0008553393,0.007730814,0.0009136648,0.0012820268,0.0023568366,0.0110999895,0.38932264],"study_design_scores_gemma":[0.0025750098,0.0038736886,0.2662761,0.36936432,0.20261788,0.005078478,0.030946737,0.006491837,0.0042654346,0.011661583,0.09594691,0.0009019846],"about_ca_topic_score_codex":0.0037500479,"about_ca_topic_score_gemma":0.008478775,"teacher_disagreement_score":0.86177516,"about_ca_system_score_codex":0.0065703094,"about_ca_system_score_gemma":0.0238706,"threshold_uncertainty_score":0.37447113},"labels":[{"model":"gemma","categories":["bibliometrics"],"domain":null,"study_design":"not_applicable","genre":"review","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":["bibliometrics"],"domain":null,"study_design":"systematic_review","genre":"review","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W4403808708","doi":"10.48550/arxiv.2409.00222","title":"Can Large Language Models Address Open-Target Stance Detection?","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Computer science; Artificial intelligence","score_opus":0.06989577715643526,"score_gpt":0.21345058819809704,"score_spread":0.1435548110416618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403808708","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16225868,0.00818636,0.77480525,0.007587408,0.0010934196,0.0004706582,0.006942103,0.026909534,0.011746678],"genre_scores_gemma":[0.7394334,0.0017161125,0.23914431,0.0016121976,0.0006034413,0.00045624224,0.010537506,0.001914986,0.004581798],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973579,0.0015318756,0.00015832232,0.00055756833,0.00024815297,0.0001461087],"domain_scores_gemma":[0.98511654,0.011580904,0.000498046,0.0014677295,0.0009829799,0.00035391664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049500167,0.0020230457,0.0012046049,0.0013703202,0.00059181184,0.0027000604,0.0017560783,0.0020686106,0.004213959],"category_scores_gemma":[0.025039172,0.0007395229,0.0012534042,0.0011110456,0.000647088,0.0068515735,0.0016221668,0.0028574823,0.0075046043],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019130787,0.00078318536,0.01800735,0.002038187,0.0007818775,0.00072068995,0.0015349064,0.18152764,0.02748298,0.015847487,0.066815116,0.68254745],"study_design_scores_gemma":[0.0001387405,0.00018151395,0.001172366,0.00011000793,0.00010376864,0.00023376173,0.0002881033,0.95046645,0.0051657036,0.030961016,0.011134661,0.000043826123],"about_ca_topic_score_codex":0.0034776654,"about_ca_topic_score_gemma":0.006943993,"teacher_disagreement_score":0.0049500167,"about_ca_system_score_codex":0.0008562906,"about_ca_system_score_gemma":0.0013774498,"threshold_uncertainty_score":0.02617848},"labels":[],"label_agreement":null},{"id":"W4403813734","doi":"10.32388/d1mvb5","title":"LFOSum: Summarizing Long-form Opinions with Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"Qeios","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Linguistics; Computer science; Natural language processing; Philosophy","score_opus":0.02904424962697813,"score_gpt":0.27747785894276833,"score_spread":0.2484336093157902,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403813734","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049739126,0.0053375545,0.86694354,0.0016118074,0.00067369075,0.0006981905,0.02160592,0.04954868,0.003841535],"genre_scores_gemma":[0.26389065,0.0017969005,0.65669066,0.0007577916,0.00070851657,0.0009656459,0.06730704,0.0015568372,0.0063260254],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982381,0.00081586663,0.00018864768,0.00037503566,0.00030212037,0.000080228114],"domain_scores_gemma":[0.9924643,0.004366578,0.00055964116,0.0010836121,0.0013421592,0.00018380018],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029420056,0.002522883,0.0008468457,0.0028445958,0.00048704955,0.0018419148,0.0014123975,0.0014748563,0.0028867056],"category_scores_gemma":[0.018138213,0.00036093296,0.0013162049,0.0017503368,0.00029759592,0.0028771032,0.0015543663,0.0016958403,0.003719851],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067050895,0.0003363441,0.005601546,0.0019327396,0.0006664677,0.00036759817,0.0010527557,0.05920407,0.018811574,0.0038786326,0.09594864,0.8115292],"study_design_scores_gemma":[0.0001574262,0.0004868407,0.0035890862,0.00016178528,0.0003366749,0.00022046085,0.0005040178,0.9142117,0.021106178,0.014324801,0.044781603,0.00011945806],"about_ca_topic_score_codex":0.0036756457,"about_ca_topic_score_gemma":0.009311983,"teacher_disagreement_score":0.0036756457,"about_ca_system_score_codex":0.00071582774,"about_ca_system_score_gemma":0.0009875147,"threshold_uncertainty_score":0.015559018},"labels":[],"label_agreement":null},{"id":"W4403881534","doi":"10.1111/1755-6724.15213","title":"GeoNER: Geological Named Entity Recognition with Enriched Domain Pre‐Training Model and Adversarial Training","year":2024,"lang":"en","type":"article","venue":"Acta Geologica Sinica - English Edition","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China; Ministry of Natural Resources","keywords":"Training (meteorology); Adversarial system; Computer science; Domain (mathematical analysis); Training set; Artificial intelligence; Named-entity recognition; Pattern recognition (psychology); Engineering; Mathematics; Geography","score_opus":0.04337033694538565,"score_gpt":0.24662427379741245,"score_spread":0.2032539368520268,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403881534","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05614119,0.0009917337,0.928572,0.00074826693,0.0002394406,0.00011460187,0.0007761087,0.0078777755,0.0045388923],"genre_scores_gemma":[0.795134,0.00048311573,0.18379912,0.00072045875,0.00014495917,0.00022290555,0.0038559975,0.0002118879,0.015427632],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99958926,0.00010710237,0.000019219106,0.0001586944,0.000065494845,0.00006017727],"domain_scores_gemma":[0.9993753,0.000299881,0.000040817682,0.00014827811,0.00010325471,0.000032420583],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00096463173,0.0010024342,0.00067904283,0.00051339687,0.0002634294,0.0006258109,0.0018033157,0.0011543103,0.0030208863],"category_scores_gemma":[0.0016066303,0.00029140033,0.0006619106,0.00054070895,0.00044343967,0.0013867089,0.0014217984,0.0017458999,0.0014194618],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002500636,0.00017300517,0.0015592697,0.00009299266,0.00009881765,0.00023779865,0.000070325266,0.6887502,0.005918269,0.0065506087,0.01119305,0.28510562],"study_design_scores_gemma":[0.0000025471868,0.000017299277,0.00009502297,0.0000032429818,0.000004701088,0.000016986956,0.0000046006035,0.9969689,0.0013556066,0.0009938668,0.00053328404,0.0000038989074],"about_ca_topic_score_codex":0.0061045825,"about_ca_topic_score_gemma":0.0063382816,"teacher_disagreement_score":0.0061045825,"about_ca_system_score_codex":0.0005314157,"about_ca_system_score_gemma":0.00066135067,"threshold_uncertainty_score":0.012138128},"labels":[],"label_agreement":null},{"id":"W4403935722","doi":"10.1145/3652620.3687807","title":"Multi-step Iterative Automated Domain Modeling with Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Mitacs","keywords":"Computer science; Domain (mathematical analysis); Modeling language; Programming language; Software; Mathematics","score_opus":0.024221829496532888,"score_gpt":0.28245915209019706,"score_spread":0.25823732259366416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403935722","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071256123,0.00024126175,0.97550833,0.00025121562,0.00002259914,0.00034162874,0.00053771,0.014856612,0.0011150159],"genre_scores_gemma":[0.056331657,0.00020620426,0.9371024,0.00021361477,0.000015374057,0.00044147635,0.0033858044,0.0008075561,0.00149586],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958621,0.0015759604,0.00032320715,0.00089988834,0.0011577363,0.00018096283],"domain_scores_gemma":[0.9922059,0.0045162756,0.00042231302,0.001761839,0.00091792695,0.00017572807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031948918,0.0022180607,0.0012499819,0.0029668575,0.0009397138,0.0027798896,0.0036311583,0.0013822016,0.0043957136],"category_scores_gemma":[0.010883308,0.0011049022,0.0039044134,0.0018766964,0.0008175833,0.004381361,0.004548786,0.002881444,0.0031150097],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029415032,0.0007178242,0.004136529,0.0011996024,0.00032169375,0.00091960514,0.0018825163,0.12415088,0.025063364,0.026619531,0.021711731,0.79298264],"study_design_scores_gemma":[0.00007218474,0.00008511778,0.00054430094,0.00007479493,0.000082779,0.00032141656,0.00042263776,0.9359236,0.013677254,0.027794402,0.020947225,0.000054257875],"about_ca_topic_score_codex":0.0068278145,"about_ca_topic_score_gemma":0.015675586,"teacher_disagreement_score":0.0068278145,"about_ca_system_score_codex":0.0017425781,"about_ca_system_score_gemma":0.0032992552,"threshold_uncertainty_score":0.016896427},"labels":[],"label_agreement":null},{"id":"W4403940973","doi":"10.1007/978-3-031-73226-3_16","title":"Model Breadcrumbs: Scaling Multi-task Model Merging with Sparse Masks","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Computer science; Scaling; Task (project management); Algorithm; Artificial intelligence; Theoretical computer science; Mathematics; Geometry","score_opus":0.03907107415938375,"score_gpt":0.2588326428358412,"score_spread":0.21976156867645744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403940973","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008108763,0.00037362453,0.9815756,0.0001872631,0.00018689413,0.0000932839,0.00023680039,0.007454699,0.001783074],"genre_scores_gemma":[0.16653103,0.00038901897,0.8213387,0.000464925,0.00020755334,0.00025801305,0.0018610511,0.0029663637,0.0059833317],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99913484,0.00017782122,0.0000491928,0.00026707136,0.00025559872,0.00011554319],"domain_scores_gemma":[0.99839824,0.00053743605,0.00007509822,0.00062524737,0.00024752828,0.00011645659],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015280567,0.0020031703,0.0021502862,0.0010046685,0.00090345106,0.002010611,0.0031222086,0.002045725,0.011145952],"category_scores_gemma":[0.005032686,0.0014319903,0.0021007403,0.0016921969,0.0006345602,0.003264433,0.0043128976,0.0030687952,0.0043853396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008443292,0.00029101365,0.0006942491,0.00025857636,0.0004090092,0.0002226862,0.00030744763,0.19259119,0.02995463,0.008459045,0.026133873,0.739834],"study_design_scores_gemma":[0.00003251174,0.000061146246,0.00015669882,0.0000121322755,0.00004185301,0.000059555856,0.000038730086,0.9783508,0.0061740945,0.011008718,0.0040425165,0.000021262449],"about_ca_topic_score_codex":0.013423682,"about_ca_topic_score_gemma":0.019925417,"teacher_disagreement_score":0.013423682,"about_ca_system_score_codex":0.00080948626,"about_ca_system_score_gemma":0.0016304478,"threshold_uncertainty_score":0.037286937},"labels":[],"label_agreement":null},{"id":"W4404002909","doi":"10.1145/3702980","title":"Trained without My Consent: Detecting Code Inclusion in Language Models Trained on Code","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"Fonds de recherche du Québec; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Code (set theory); Inclusion (mineral); Programming language; Physics","score_opus":0.10487511590320443,"score_gpt":0.33703465410834965,"score_spread":0.23215953820514523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404002909","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3016911,0.0019015402,0.58236516,0.007719772,0.0010717622,0.0014420899,0.033166267,0.05346269,0.017179696],"genre_scores_gemma":[0.7038024,0.00032296023,0.22867717,0.002722151,0.00020807619,0.0012406096,0.05291411,0.0016638294,0.008448721],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912145,0.0033542456,0.00062508095,0.0023789778,0.0019524347,0.00047482506],"domain_scores_gemma":[0.9630858,0.018177656,0.002245161,0.011643349,0.004089282,0.0007587372],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010788836,0.0009933998,0.00077092217,0.0012254332,0.00075715437,0.001926035,0.0018653753,0.0016700004,0.004654708],"category_scores_gemma":[0.068947665,0.0005505059,0.00096285035,0.00083342445,0.0011151519,0.00312864,0.0030705126,0.002805435,0.004878815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001762024,0.0006822518,0.09749233,0.000804162,0.00036208282,0.0018392222,0.0019930915,0.041946776,0.019336676,0.010288839,0.13154899,0.6919435],"study_design_scores_gemma":[0.00024005598,0.00041250812,0.018521024,0.00042203302,0.00012640569,0.0009485592,0.0006448225,0.85644585,0.026458694,0.037392415,0.058231685,0.00015592999],"about_ca_topic_score_codex":0.0066916468,"about_ca_topic_score_gemma":0.013144397,"teacher_disagreement_score":0.010788836,"about_ca_system_score_codex":0.00094107655,"about_ca_system_score_gemma":0.0036243794,"threshold_uncertainty_score":0.05705756},"labels":[],"label_agreement":null},{"id":"W4404008801","doi":"10.1007/978-981-97-8487-5_23","title":"AtomTool: Empowering Large Language Models with Tool Utilization Skills","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Geomechanica (Canada)","funders":"","keywords":"Computer science; Software engineering; Programming language; Human–computer interaction","score_opus":0.017409433068775453,"score_gpt":0.2632349959110514,"score_spread":0.2458255628422759,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404008801","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0074336343,0.00016014265,0.88796425,0.0002931001,0.00009131878,0.00012287736,0.0016754526,0.094465755,0.007793391],"genre_scores_gemma":[0.15173928,0.000747862,0.7785426,0.00047388647,0.000083925595,0.0005482426,0.012015031,0.033328537,0.022520617],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991947,0.00018267111,0.0000626971,0.00016301611,0.00033447333,0.0000625144],"domain_scores_gemma":[0.997615,0.0014784102,0.0000787389,0.0005466858,0.00019518475,0.00008607066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010878863,0.0014547312,0.00060361385,0.0009423268,0.0004939898,0.002569581,0.0022969544,0.0009474904,0.020019634],"category_scores_gemma":[0.0054916646,0.0014542416,0.0016781393,0.0008131833,0.0005717703,0.005423681,0.0036578197,0.002318137,0.0100583825],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066890137,0.00042329438,0.0041471543,0.0013135816,0.00025111108,0.0008446911,0.0018633463,0.047364652,0.035826977,0.08290894,0.1680979,0.6562894],"study_design_scores_gemma":[0.00013904016,0.000122049416,0.00052132184,0.00018691119,0.00018863675,0.0005772234,0.00037251046,0.59647965,0.04202912,0.08951729,0.26975948,0.00010679913],"about_ca_topic_score_codex":0.0021816844,"about_ca_topic_score_gemma":0.00493336,"teacher_disagreement_score":0.020019634,"about_ca_system_score_codex":0.0004904914,"about_ca_system_score_gemma":0.0010914109,"threshold_uncertainty_score":0.066972315},"labels":[],"label_agreement":null},{"id":"W4404021177","doi":"10.1007/s44163-024-00175-8","title":"A survey on augmenting knowledge graphs (KGs) with large language models (LLMs): models, evaluation metrics, benchmarks, and challenges","year":2024,"lang":"en","type":"article","venue":"Discover Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Toronto Metropolitan University","funders":"","keywords":"Computer science; Data science; Natural language processing","score_opus":0.15560490023979645,"score_gpt":0.3427390326293876,"score_spread":0.18713413238959115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404021177","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064327054,0.31326532,0.5630003,0.009487691,0.0006788916,0.0009419128,0.0055867885,0.010339272,0.03237273],"genre_scores_gemma":[0.2393884,0.14358632,0.59955186,0.0016657211,0.0006638493,0.0007823127,0.010802914,0.0015995604,0.0019591008],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98057437,0.009674644,0.00169193,0.001884948,0.005690064,0.00048405316],"domain_scores_gemma":[0.86219066,0.11559565,0.0031095804,0.009331219,0.008551327,0.001221606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021390878,0.0028261247,0.0023273716,0.011151604,0.00075725536,0.0065522543,0.0034800293,0.0020273526,0.0028332828],"category_scores_gemma":[0.10179394,0.0011597759,0.0019059443,0.013516842,0.0015255422,0.012019678,0.0034988902,0.0025006211,0.0013293668],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034031257,0.0004427538,0.012793186,0.008273667,0.0006381223,0.00008084342,0.00072124746,0.057378843,0.0016965625,0.03757945,0.02109951,0.85895544],"study_design_scores_gemma":[0.00010820512,0.0012840049,0.015337987,0.009376661,0.0012823694,0.0007134819,0.0022720299,0.62136155,0.010138562,0.16914387,0.16861942,0.00036183235],"about_ca_topic_score_codex":0.012636872,"about_ca_topic_score_gemma":0.010952576,"teacher_disagreement_score":0.021390878,"about_ca_system_score_codex":0.0034702634,"about_ca_system_score_gemma":0.0040687597,"threshold_uncertainty_score":0.11312711},"labels":[],"label_agreement":null},{"id":"W4404280427","doi":"10.23977/acss.2024.080618","title":"Sparse Attention Mechanisms in Large Language Models: Applications, Classification, Performance Analysis, and Optimization","year":2024,"lang":"en","type":"article","venue":"Advances in Computer Signals and Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Machine learning; Natural language processing; Pattern recognition (psychology)","score_opus":0.02042878749782472,"score_gpt":0.26761940755077795,"score_spread":0.24719062005295322,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404280427","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026615186,0.0018305926,0.96716624,0.0010768265,0.00005520925,0.0000751206,0.00010873995,0.0012140512,0.001858043],"genre_scores_gemma":[0.6601472,0.0026637123,0.33058536,0.0004890617,0.0002970319,0.0003402564,0.00057946396,0.00031320177,0.0045847003],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986168,0.0006766583,0.00006900618,0.00022891743,0.00026539466,0.0001431832],"domain_scores_gemma":[0.9926416,0.0060261725,0.00029296946,0.00044578774,0.0004671363,0.00012627948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043269945,0.001311305,0.0015045353,0.0011032579,0.00057329604,0.0018259727,0.0016458279,0.0014599299,0.0026270873],"category_scores_gemma":[0.015221586,0.0005434541,0.0008396019,0.0014669555,0.0010555619,0.0041501834,0.001812294,0.0021845824,0.00082828547],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025080077,0.00019328701,0.0015352997,0.0002850414,0.00013763919,0.00006822737,0.00019799148,0.6553121,0.0046036965,0.046465676,0.0040727565,0.2868774],"study_design_scores_gemma":[0.0000061552396,0.000029109877,0.00010458894,0.0000066270327,0.000010859324,0.000011389334,0.000013459032,0.9866284,0.00079428137,0.012089632,0.00029913566,0.000006278269],"about_ca_topic_score_codex":0.011438298,"about_ca_topic_score_gemma":0.010504827,"teacher_disagreement_score":0.011438298,"about_ca_system_score_codex":0.0018018897,"about_ca_system_score_gemma":0.0015873241,"threshold_uncertainty_score":0.022883654},"labels":[],"label_agreement":null},{"id":"W4404356684","doi":"10.1007/978-3-031-73503-5_4","title":"AICIS: A System for Identifying AI Contribution in Textual Content","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Content (measure theory); Information retrieval; Natural language processing; Artificial intelligence; Mathematics","score_opus":0.04707489017599698,"score_gpt":0.28308892216841497,"score_spread":0.236014031992418,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404356684","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05044241,0.0015086007,0.448405,0.00051357516,0.00050223974,0.0012852666,0.07228595,0.38988438,0.035172656],"genre_scores_gemma":[0.15073885,0.00072936463,0.7198607,0.00028545476,0.0004384562,0.0018411408,0.091277644,0.0077735675,0.027054796],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99861956,0.00020353735,0.00013205934,0.00039785515,0.0005686851,0.00007840563],"domain_scores_gemma":[0.993436,0.0033683563,0.0005081627,0.0006914629,0.0016140925,0.00038204473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00173827,0.0015798256,0.00083244266,0.01301538,0.0010089691,0.0025504017,0.001304062,0.0010033203,0.022305094],"category_scores_gemma":[0.010602213,0.00048669727,0.00071624224,0.006763393,0.0003981286,0.0036481665,0.0023772924,0.0009766737,0.016354106],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011524914,0.00021803667,0.020742511,0.0013276422,0.00021235306,0.00029716443,0.0013227897,0.0014372043,0.021432037,0.0058827396,0.17236048,0.7736146],"study_design_scores_gemma":[0.00026018507,0.00057522947,0.055024084,0.0005068353,0.0008730336,0.0011801564,0.0023907674,0.39925516,0.08973622,0.023149418,0.42676336,0.00028558055],"about_ca_topic_score_codex":0.006051014,"about_ca_topic_score_gemma":0.010734519,"teacher_disagreement_score":0.022305094,"about_ca_system_score_codex":0.0010076247,"about_ca_system_score_gemma":0.0014928351,"threshold_uncertainty_score":0.07461798},"labels":[],"label_agreement":null},{"id":"W4404366193","doi":"10.1007/s10115-024-02269-2","title":"An evidence-based approach for open-domain question answering","year":2024,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Question answering; Information retrieval; Relevance (law); Open domain; Benchmark (surveying); Context (archaeology); Graph; Construct (python library); Domain (mathematical analysis); Rank (graph theory); Knowledge graph; Artificial intelligence; Natural language processing; Theoretical computer science","score_opus":0.05172166446128868,"score_gpt":0.3110074057323632,"score_spread":0.2592857412710745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404366193","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003533508,0.00069479825,0.99125683,0.0012074136,0.00006746026,0.00020945593,0.00050776254,0.00070343097,0.0018193129],"genre_scores_gemma":[0.14353503,0.0006039762,0.8507194,0.00038307987,0.0002351769,0.0004095359,0.0018820334,0.00012907085,0.0021026877],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99002516,0.0042963447,0.0011576712,0.0015503146,0.0026567655,0.0003137636],"domain_scores_gemma":[0.9604143,0.03180639,0.001133372,0.0023874668,0.0035804948,0.0006781162],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0088766515,0.0011675428,0.0021584837,0.008564285,0.0017152188,0.0054906723,0.0048756907,0.0044106846,0.0077177626],"category_scores_gemma":[0.046794035,0.0011998102,0.0030599614,0.0059729707,0.0018357543,0.008595634,0.0057880343,0.0043659937,0.0021195726],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007208426,0.0009017002,0.0047578067,0.0014387921,0.00071806245,0.00078599993,0.0018627165,0.04780253,0.007245023,0.19426104,0.013748316,0.7257572],"study_design_scores_gemma":[0.00013399987,0.00012212877,0.0012800212,0.00032268482,0.00042201197,0.00044572944,0.00049679185,0.6428141,0.0044425814,0.33252805,0.016899427,0.00009245564],"about_ca_topic_score_codex":0.0051901164,"about_ca_topic_score_gemma":0.008744657,"teacher_disagreement_score":0.0088766515,"about_ca_system_score_codex":0.0017311564,"about_ca_system_score_gemma":0.0030967821,"threshold_uncertainty_score":0.046944797},"labels":[],"label_agreement":null},{"id":"W4404395182","doi":"10.2196/60272","title":"Enhancing Bias Assessment for Complex Term Groups in Language Embedding Models: Quantitative Comparison of Methods","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Word embedding; Measure (data warehouse); Word2vec; Machine learning; Embedding; Association test; Term (time); Data mining","score_opus":0.17317881202963353,"score_gpt":0.5031536617832251,"score_spread":0.32997484975359154,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404395182","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29214153,0.008765895,0.6889445,0.0008999521,0.0006379826,0.0013321768,0.0016924449,0.0012725972,0.0043129656],"genre_scores_gemma":[0.7033762,0.0009380502,0.290898,0.00022705083,0.00019771248,0.0013377737,0.002051606,0.00035754158,0.0006160708],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95781523,0.030416157,0.0030452525,0.0031591926,0.0050806003,0.00048348762],"domain_scores_gemma":[0.5627746,0.39543417,0.010018175,0.014570285,0.015958969,0.001243784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08726801,0.0018472264,0.0011427875,0.00410594,0.0007225629,0.0026760648,0.001254331,0.0019855003,0.0024904876],"category_scores_gemma":[0.24793802,0.00043937797,0.0023014182,0.0022998925,0.0016659026,0.0045251916,0.0037493373,0.0022670773,0.0005515992],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0075502307,0.0011084131,0.12973452,0.00400575,0.0061440393,0.00016424121,0.0024803882,0.17684737,0.0064453483,0.01784906,0.0073612193,0.64030945],"study_design_scores_gemma":[0.0007162102,0.0028477996,0.04078424,0.0011896946,0.0015956094,0.0003417591,0.0013106926,0.8883845,0.01300752,0.04345946,0.006019703,0.00034286396],"about_ca_topic_score_codex":0.0019054852,"about_ca_topic_score_gemma":0.002099666,"teacher_disagreement_score":0.08726801,"about_ca_system_score_codex":0.0014470773,"about_ca_system_score_gemma":0.0017963912,"threshold_uncertainty_score":0.46152288},"labels":[],"label_agreement":null},{"id":"W4404412796","doi":"10.1016/j.eswa.2024.125648","title":"A new approach for competency frameworks mapping using large language models","year":2024,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Université TÉLUQ","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Data science","score_opus":0.03832938816186795,"score_gpt":0.2950653646710826,"score_spread":0.25673597650921465,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404412796","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00095338793,0.0000728928,0.995808,0.00017141829,0.000043977154,0.00006221149,0.00025652567,0.0014583921,0.0011731558],"genre_scores_gemma":[0.055318624,0.00022378587,0.9384555,0.00019928467,0.000071940834,0.0003049842,0.0015842145,0.00065702107,0.0031846361],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99498796,0.0014774202,0.00044348484,0.001164479,0.001715283,0.00021148149],"domain_scores_gemma":[0.99399513,0.0025727395,0.00023782156,0.001668079,0.001253102,0.0002731895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031378292,0.001155357,0.0011701448,0.0038794263,0.0017308927,0.0044923113,0.0030366755,0.0015279134,0.007336215],"category_scores_gemma":[0.013140687,0.0011817586,0.003711237,0.0030402204,0.0010437354,0.008628591,0.006479226,0.004498946,0.0036540802],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019168165,0.00042497367,0.0024762903,0.0006040387,0.0004301336,0.000533591,0.0021766257,0.04163129,0.011176356,0.33505777,0.02097641,0.5843208],"study_design_scores_gemma":[0.00003539387,0.00006770779,0.0007414659,0.00014910246,0.00015437456,0.0005038845,0.0006736601,0.6012244,0.006823597,0.33466065,0.054855987,0.000109795415],"about_ca_topic_score_codex":0.008542319,"about_ca_topic_score_gemma":0.014066769,"teacher_disagreement_score":0.008542319,"about_ca_system_score_codex":0.0014233482,"about_ca_system_score_gemma":0.0032046272,"threshold_uncertainty_score":0.024542153},"labels":[],"label_agreement":null},{"id":"W4404609282","doi":"10.1109/tse.2024.3504286","title":"On Inter-Dataset Code Duplication and Data Leakage in Large Language Models","year":2024,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; McGill University","funders":"","keywords":"Computer science; Programming language; Code (set theory); Software bug; Data mining; Software","score_opus":0.025993329642164737,"score_gpt":0.2756331319817165,"score_spread":0.24963980233955177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404609282","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3296938,0.004565615,0.6286742,0.01354997,0.0005207661,0.00065116375,0.0047483374,0.012854946,0.0047411886],"genre_scores_gemma":[0.833484,0.0008029927,0.14889881,0.0029825754,0.00038314963,0.0006379576,0.006726351,0.0015002065,0.0045840065],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9706745,0.013999949,0.002330345,0.006844714,0.004346288,0.0018041177],"domain_scores_gemma":[0.76979864,0.16939682,0.009390376,0.043318123,0.006304268,0.0017917984],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028724778,0.0014985979,0.002957208,0.0030881942,0.0030637549,0.0045601856,0.0057495544,0.003566684,0.002535812],"category_scores_gemma":[0.15623026,0.001882936,0.002995149,0.00584463,0.0040215175,0.013612632,0.008428127,0.0044343504,0.0010557726],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023954837,0.00091094064,0.059938576,0.00080977794,0.0010049612,0.0015263899,0.0019723775,0.4896173,0.0044682967,0.058665648,0.037367295,0.34132293],"study_design_scores_gemma":[0.00010692558,0.00015567629,0.0025127959,0.00008018933,0.00013072017,0.0004471523,0.00026525228,0.91814876,0.003195175,0.0723694,0.002545991,0.000042081992],"about_ca_topic_score_codex":0.011486677,"about_ca_topic_score_gemma":0.014108108,"teacher_disagreement_score":0.9712752,"about_ca_system_score_codex":0.006384215,"about_ca_system_score_gemma":0.006746542,"threshold_uncertainty_score":0.15191293},"labels":[],"label_agreement":null},{"id":"W4404729236","doi":"10.1111/exsy.13789","title":"Flexible Distribution Approaches to Enhance Regression and Deep Topic Modelling Techniques","year":2024,"lang":"en","type":"article","venue":"Expert Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Regression; Distribution (mathematics); Regression analysis; Artificial intelligence; Data mining; Data science; Machine learning; Statistics; Mathematics","score_opus":0.09131876294557813,"score_gpt":0.29966527382244984,"score_spread":0.20834651087687173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404729236","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002904193,0.0002389038,0.99587655,0.00010280201,0.000018728959,0.00002081578,0.0000551104,0.00046863453,0.00031423502],"genre_scores_gemma":[0.33094472,0.0012605457,0.65907943,0.00041314674,0.00034997863,0.0004932737,0.0012920467,0.0008767107,0.0052901423],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99722606,0.0015335791,0.00012793721,0.00055975403,0.0003778358,0.00017490878],"domain_scores_gemma":[0.99416625,0.004429828,0.00028452833,0.0004896001,0.0005094897,0.00012031039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055408245,0.0015433436,0.0017804406,0.0026043777,0.00057062996,0.0016287753,0.0031118086,0.0017227198,0.0033921117],"category_scores_gemma":[0.014565745,0.00095843925,0.0022865694,0.0028452384,0.0009983815,0.0036202178,0.002752258,0.0035300995,0.002058815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016750206,0.00015484834,0.0018264181,0.00021485919,0.0002148177,0.0001634056,0.00046498896,0.6393737,0.005121193,0.07631348,0.004693827,0.27129096],"study_design_scores_gemma":[0.000008014634,0.000008817755,0.00009050907,0.000008517502,0.000008901554,0.00001626373,0.000009607678,0.98116356,0.00041955532,0.017390816,0.00086743076,0.000007921397],"about_ca_topic_score_codex":0.005113065,"about_ca_topic_score_gemma":0.0064701433,"teacher_disagreement_score":0.0055408245,"about_ca_system_score_codex":0.0011662157,"about_ca_system_score_gemma":0.001090526,"threshold_uncertainty_score":0.029303014},"labels":[],"label_agreement":null},{"id":"W4404764563","doi":"10.1038/s41562-024-02046-9","title":"Large language models surpass human experts in predicting neuroscience results","year":2024,"lang":"en","type":"article","venue":"Nature Human Behaviour","topic":"Topic Modeling","field":"Computer Science","cited_by":111,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut Universitaire de Gériatrie de Montréal","funders":"National Eye Institute; Economic and Social Research Council; Research Councils UK; Royal Society","keywords":"Cognitive science; Psychology; Neuroscience","score_opus":0.029284502508189935,"score_gpt":0.3250620535175684,"score_spread":0.29577755100937847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404764563","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5938138,0.007816267,0.34872854,0.0108250575,0.00096946297,0.00026253462,0.005344354,0.011757453,0.020482562],"genre_scores_gemma":[0.9283525,0.0006596078,0.06254801,0.0013155513,0.00028166894,0.000120563294,0.0039196312,0.000482344,0.0023201993],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943638,0.0031253074,0.0003271175,0.0012477563,0.0006588225,0.00027709996],"domain_scores_gemma":[0.93847,0.05180813,0.002286199,0.0038025198,0.0019703354,0.0016628478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015029472,0.001968106,0.0009912736,0.0023155008,0.000576195,0.0033866388,0.0009935427,0.0027835593,0.0030605376],"category_scores_gemma":[0.06061502,0.0005306891,0.0010746077,0.00085591775,0.0010469964,0.0059339264,0.0017279859,0.0029543373,0.0020644993],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029441814,0.0007395688,0.09122994,0.0015196052,0.0013315685,0.00074391847,0.0020611614,0.47430485,0.01310379,0.01724701,0.04985119,0.3449232],"study_design_scores_gemma":[0.000113735434,0.00037769953,0.007052858,0.00013237185,0.00013408353,0.00016550336,0.0002704205,0.93221456,0.005508805,0.04729022,0.006641774,0.00009804636],"about_ca_topic_score_codex":0.004821627,"about_ca_topic_score_gemma":0.008664686,"teacher_disagreement_score":0.015029472,"about_ca_system_score_codex":0.0010250879,"about_ca_system_score_gemma":0.0017535132,"threshold_uncertainty_score":0.0794844},"labels":[],"label_agreement":null},{"id":"W4404774079","doi":"10.1101/2024.11.27.625607","title":"Rapid Semantic Processing: An MEG Study of Narrative Text Reading","year":2024,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Optech (Canada)","funders":"","keywords":"Reading (process); Narrative; Natural language processing; Computer science; Linguistics; Artificial intelligence; Psychology; Philosophy","score_opus":0.02707372521275579,"score_gpt":0.25577564852538204,"score_spread":0.22870192331262626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404774079","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99676275,0.000098722456,0.0017649767,0.00005951674,0.0000071562617,0.000033958673,0.00015416341,0.000024251794,0.0010945409],"genre_scores_gemma":[0.9969689,0.00010690365,0.0018782899,0.00005275833,0.00003040319,0.00004204496,0.00017058133,0.000019563224,0.0007305361],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.99993885,0.000017268947,0.0000036486524,0.000019663968,0.000013037697,0.0000074864624],"domain_scores_gemma":[0.99979407,0.00012421461,0.000032904838,0.000018373796,0.000015189968,0.000015380003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00012581756,0.00014355975,0.00018824398,0.00023034569,0.0001427051,0.00028959705,0.00014019494,0.0003019382,0.0012700217],"category_scores_gemma":[0.00121719,0.000096486605,0.00012622819,0.00025236473,0.0003656423,0.00025431154,0.0001933477,0.00026391205,0.0002995298],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022331025,0.0004241648,0.029410452,0.00023888574,0.000085433276,0.004110137,0.004876737,0.00028596117,0.9081326,0.0011177284,0.0008151762,0.048269574],"study_design_scores_gemma":[0.0002822292,0.0022697132,0.9026432,0.000035689074,0.00010554818,0.011079905,0.0034436402,0.0045360588,0.064896904,0.00481902,0.005836289,0.00005188245],"about_ca_topic_score_codex":0.00042530583,"about_ca_topic_score_gemma":0.00068914506,"teacher_disagreement_score":0.0012700217,"about_ca_system_score_codex":0.000093252405,"about_ca_system_score_gemma":0.000071481765,"threshold_uncertainty_score":0.0042486787},"labels":[],"label_agreement":null},{"id":"W4404780711","doi":"10.18653/v1/2024.wnu-1.4","title":"Using Large Language Models for Understanding Narrative Discourse","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Narrative; Computer science; Linguistics; Natural language processing; Philosophy","score_opus":0.149778524103757,"score_gpt":0.3744670122599409,"score_spread":0.2246884881561839,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404780711","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03859578,0.0014603363,0.94867176,0.0017506236,0.000106099455,0.00031135775,0.00247371,0.0028948085,0.0037355719],"genre_scores_gemma":[0.48096684,0.00089530187,0.5074268,0.00032185836,0.0001474008,0.0011719689,0.0063060005,0.00058976986,0.0021740994],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976164,0.00166377,0.00011196642,0.00039569102,0.00016985518,0.000042405074],"domain_scores_gemma":[0.9849902,0.013016762,0.00063824275,0.00074986776,0.00045308372,0.00015179616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040504504,0.0012447508,0.00055860664,0.0023602948,0.00091671676,0.0037257418,0.001913712,0.0012050638,0.0028436764],"category_scores_gemma":[0.023952926,0.00066446536,0.0015554529,0.0013380655,0.0009163666,0.0061333897,0.0019449124,0.002195935,0.0011568101],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005968059,0.0003306564,0.013274452,0.0019430796,0.00072992913,0.0012122368,0.020310728,0.47084355,0.0099818595,0.18813737,0.016113799,0.27652556],"study_design_scores_gemma":[0.00003018789,0.000026942209,0.00056427513,0.00010617456,0.000050221035,0.00010365709,0.0008830831,0.9045454,0.0012079597,0.079515845,0.012938454,0.000027751199],"about_ca_topic_score_codex":0.004999974,"about_ca_topic_score_gemma":0.010186467,"teacher_disagreement_score":0.004999974,"about_ca_system_score_codex":0.0018976299,"about_ca_system_score_gemma":0.0011888598,"threshold_uncertainty_score":0.021421075},"labels":[],"label_agreement":null},{"id":"W4404781236","doi":"10.18653/v1/2024.nllp-1.5","title":"Quebec Automobile Insurance Question-Answering With Retrieval-Augmented Generation","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université Laval","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada","keywords":"Question answering; Automobile insurance; Computer science; Information retrieval; Business; Actuarial science","score_opus":0.01546751376813193,"score_gpt":0.25044122389066975,"score_spread":0.23497371012253782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404781236","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58417916,0.0048724688,0.13280621,0.0042264005,0.00058485527,0.0034590708,0.13453272,0.06481869,0.07052042],"genre_scores_gemma":[0.64273095,0.00074166723,0.13353188,0.001070614,0.00015758151,0.0013494063,0.19115172,0.0015421128,0.027723968],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973295,0.0010612279,0.00012606198,0.0007845874,0.0005136934,0.00018491126],"domain_scores_gemma":[0.9922746,0.0029674936,0.00020680892,0.0011169583,0.00310438,0.00032979148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027649454,0.0014305223,0.000594949,0.0030577949,0.0018498254,0.0014759874,0.0020073585,0.0017439206,0.017843448],"category_scores_gemma":[0.012502969,0.00042602228,0.00075160246,0.0023493394,0.00093685085,0.0018096855,0.0015700407,0.0014083325,0.005381121],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012208069,0.0009835748,0.021888776,0.0023937335,0.00025418747,0.0012494536,0.004588398,0.031654574,0.038336083,0.005312548,0.38696453,0.5051534],"study_design_scores_gemma":[0.0010316987,0.00078000606,0.10319954,0.00039739543,0.00030790846,0.0019567541,0.003973032,0.41248554,0.07348549,0.00406063,0.39793944,0.0003825414],"about_ca_topic_score_codex":0.62075543,"about_ca_topic_score_gemma":0.6940683,"teacher_disagreement_score":0.37924457,"about_ca_system_score_codex":0.006733097,"about_ca_system_score_gemma":0.0059270505,"threshold_uncertainty_score":0.762956},"labels":[],"label_agreement":null},{"id":"W4404782347","doi":"10.18653/v1/2024.emnlp-main.734","title":"FAC2E: Better Understanding Large Language Model Capabilities by Dissociating Language and Cognition","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Computer science; Cognition; Cognitive science; Natural language processing; Language model; Artificial intelligence; Psychology","score_opus":0.022957726093236014,"score_gpt":0.2646405058209287,"score_spread":0.24168277972769267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782347","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17068657,0.0006026357,0.79607993,0.0007946799,0.000057659934,0.00045279865,0.0019797918,0.018224165,0.011121692],"genre_scores_gemma":[0.67602086,0.0002127227,0.31852627,0.00019613068,0.000021027758,0.00024467826,0.0023129764,0.0005378711,0.0019274843],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998489,0.0005876666,0.00012770525,0.00029034456,0.00037486744,0.00013045808],"domain_scores_gemma":[0.9875299,0.007913351,0.00078281824,0.0025487556,0.00086523907,0.00035991956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004288774,0.0016822274,0.0006868525,0.0026537005,0.0004948652,0.0041256235,0.0014659481,0.0013620678,0.006834712],"category_scores_gemma":[0.021480076,0.00042475984,0.00091943145,0.0010718906,0.0008617259,0.008700878,0.0040577897,0.0015661042,0.0012732901],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010418034,0.00071619,0.0437098,0.0010883192,0.00044930453,0.00028329474,0.0029592272,0.112714596,0.030088872,0.0365751,0.010695001,0.7596784],"study_design_scores_gemma":[0.000058173136,0.00034352293,0.0126247415,0.00010058493,0.00014053988,0.00025193245,0.0010249517,0.8953321,0.029600853,0.049645647,0.010736501,0.0001405352],"about_ca_topic_score_codex":0.008342521,"about_ca_topic_score_gemma":0.01059152,"teacher_disagreement_score":0.008342521,"about_ca_system_score_codex":0.0010424659,"about_ca_system_score_gemma":0.0018557108,"threshold_uncertainty_score":0.022864401},"labels":[],"label_agreement":null},{"id":"W4404782398","doi":"10.18653/v1/2024.emnlp-main.606","title":"Stable Language Model Pre-training by Reducing Embedding Variability","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Computer science; Training (meteorology); Embedding; Language model; Artificial intelligence; Natural language processing","score_opus":0.023901279584408783,"score_gpt":0.2959054770229754,"score_spread":0.27200419743856663,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782398","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10092837,0.0011886619,0.877589,0.0007146661,0.00025770706,0.00025061623,0.0005706853,0.015024333,0.003475931],"genre_scores_gemma":[0.72541296,0.00049701234,0.25933602,0.0008047438,0.00016243897,0.0006125988,0.0033386352,0.0022913523,0.0075443266],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99847466,0.0004043787,0.000117221,0.00050876953,0.00027859892,0.00021626036],"domain_scores_gemma":[0.9952003,0.0024366528,0.0002679982,0.00091517816,0.0009875858,0.00019227894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026660364,0.0028661378,0.0012291418,0.00092598464,0.00090336415,0.0020571689,0.0022033877,0.0016790343,0.005580188],"category_scores_gemma":[0.016694281,0.0010717834,0.0012062927,0.0007994061,0.0008712594,0.004384035,0.0029477654,0.005285389,0.0042002727],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072642334,0.0005418722,0.0076950192,0.00039968302,0.00032584765,0.00041825426,0.00067585293,0.30915076,0.08040499,0.007440934,0.01827691,0.5739434],"study_design_scores_gemma":[0.00004172618,0.00017993205,0.0014024253,0.0000383115,0.00006979781,0.000110238616,0.00009401311,0.9546521,0.03369811,0.0073811268,0.0022965954,0.000035665966],"about_ca_topic_score_codex":0.0054760817,"about_ca_topic_score_gemma":0.01290237,"teacher_disagreement_score":0.005580188,"about_ca_system_score_codex":0.0010839641,"about_ca_system_score_gemma":0.0026416706,"threshold_uncertainty_score":0.018667579},"labels":[],"label_agreement":null},{"id":"W4404782438","doi":"10.18653/v1/2024.emnlp-main.551","title":"StablePrompt : Automatic Prompt Tuning using Reinforcement Learning for Large Language Model","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea","keywords":"Reinforcement learning; Computer science; Language model; Artificial intelligence","score_opus":0.037925025208956874,"score_gpt":0.2995942561215851,"score_spread":0.2616692309126282,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782438","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013447977,0.00050211675,0.95115536,0.000278772,0.0001833912,0.00019785209,0.00024697592,0.032697335,0.0012902839],"genre_scores_gemma":[0.39490142,0.0003062478,0.5956582,0.000628388,0.00016768515,0.00079437357,0.0012771246,0.0027296364,0.00353688],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998579,0.00062190095,0.00007562113,0.00040517104,0.00021132758,0.00010706711],"domain_scores_gemma":[0.99559754,0.0030608866,0.00020991442,0.0005313588,0.0003629515,0.00023730184],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00298895,0.0013310233,0.0012937924,0.00053915964,0.00051509164,0.0011972939,0.00288211,0.0016513215,0.0053380784],"category_scores_gemma":[0.013291337,0.00070826407,0.00071625767,0.0004942634,0.0008819927,0.0027018355,0.0026898473,0.0037355106,0.0028407855],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008969785,0.0006882786,0.0027519776,0.0006174574,0.00012336402,0.0003460324,0.0005233651,0.29690063,0.020795044,0.01138508,0.03407093,0.6309008],"study_design_scores_gemma":[0.00007649602,0.00004943605,0.000085354644,0.000009209348,0.0000073318856,0.000022930493,0.000016380247,0.99020886,0.0022413542,0.005721692,0.0015502142,0.000010706499],"about_ca_topic_score_codex":0.0027779024,"about_ca_topic_score_gemma":0.004173557,"teacher_disagreement_score":0.0053380784,"about_ca_system_score_codex":0.00081573514,"about_ca_system_score_gemma":0.0019139788,"threshold_uncertainty_score":0.017857611},"labels":[],"label_agreement":null},{"id":"W4404782650","doi":"10.18653/v1/2024.emnlp-main.579","title":"MixGR: Enhancing Retriever Generalization for Scientific Domain through Complementary Granularity","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Bundesministerium für Bildung und Forschung","keywords":"Granularity; Generalization; Computer science; Domain (mathematical analysis); Labrador Retriever; Mathematics; Medicine; Programming language","score_opus":0.04566579330471872,"score_gpt":0.2979242531118234,"score_spread":0.2522584598071047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782650","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20962112,0.0080451,0.6953347,0.0007303553,0.00022505775,0.0008842253,0.004082026,0.073414065,0.0076633925],"genre_scores_gemma":[0.4657717,0.001488127,0.5111461,0.0006881702,0.00031151457,0.0004616155,0.011704317,0.0021245342,0.0063039395],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967855,0.000788248,0.0002594742,0.00095373887,0.000979535,0.0002334414],"domain_scores_gemma":[0.994531,0.0021309457,0.0003517154,0.0021201356,0.0006445349,0.00022174724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038494477,0.0019923002,0.0023169476,0.008174028,0.0008501881,0.001967653,0.0023677521,0.0017618169,0.0031575914],"category_scores_gemma":[0.011029054,0.0005205763,0.0016072064,0.0048901252,0.001100798,0.004994338,0.0042820494,0.0013581669,0.0036016414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009930447,0.0008058872,0.0087606255,0.0012926866,0.0004495943,0.00041275864,0.001006995,0.028003883,0.08326381,0.003971917,0.025953695,0.8450852],"study_design_scores_gemma":[0.0004433077,0.0015524084,0.014238545,0.00013217726,0.00059321127,0.0024896483,0.0011532314,0.81931937,0.09422994,0.023657875,0.04187899,0.00031127714],"about_ca_topic_score_codex":0.00366089,"about_ca_topic_score_gemma":0.00564991,"teacher_disagreement_score":0.008174028,"about_ca_system_score_codex":0.00082449825,"about_ca_system_score_gemma":0.0010824656,"threshold_uncertainty_score":0.020358086},"labels":[],"label_agreement":null},{"id":"W4404782654","doi":"10.18653/v1/2024.emnlp-main.382","title":"MirrorStories: Reflecting Diversity through Personalized Narrative Generation with Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Narrative; Diversity (politics); Computer science; Natural language generation; Natural language processing; Linguistics; Sociology; Natural language; Anthropology; Philosophy","score_opus":0.0900600388191288,"score_gpt":0.3256336003900595,"score_spread":0.2355735615709307,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782654","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33664435,0.001621199,0.62604964,0.0017714303,0.00030595541,0.0011194341,0.0035517318,0.017214159,0.011722085],"genre_scores_gemma":[0.6026342,0.00034373667,0.3856315,0.0003294803,0.00009552445,0.00071490044,0.005378526,0.00084280135,0.0040293],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99604046,0.002907685,0.00013710884,0.00047989367,0.00036179097,0.000073069714],"domain_scores_gemma":[0.9854989,0.011174099,0.0006088218,0.0018726153,0.000615606,0.00023001457],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051415428,0.0012034911,0.00048792156,0.0010836341,0.0005789812,0.0023712239,0.0015404746,0.001053903,0.004592804],"category_scores_gemma":[0.024875814,0.00047260098,0.00076976407,0.0005710302,0.0007484219,0.0040032296,0.002745907,0.0012204807,0.0017232775],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018251945,0.0008494801,0.0190001,0.0022869413,0.0005171448,0.0011953672,0.02128681,0.105656646,0.043116216,0.023471195,0.038844347,0.7419506],"study_design_scores_gemma":[0.0003194993,0.00067463866,0.0049641943,0.00026535825,0.00023980757,0.0009362658,0.005448091,0.8417237,0.03900497,0.03132729,0.07489292,0.00020327822],"about_ca_topic_score_codex":0.001429424,"about_ca_topic_score_gemma":0.0036421919,"teacher_disagreement_score":0.0051415428,"about_ca_system_score_codex":0.00067740533,"about_ca_system_score_gemma":0.00068915,"threshold_uncertainty_score":0.0271914},"labels":[],"label_agreement":null},{"id":"W4404782721","doi":"10.18653/v1/2024.emnlp-main.373","title":"Unifying Multimodal Retrieval via Document Screenshot Embedding","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Embedding; Document retrieval; Artificial intelligence; Natural language processing","score_opus":0.02529742333990012,"score_gpt":0.3007426239319106,"score_spread":0.2754452005920105,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06496831,0.0047922083,0.8893897,0.00076772866,0.00043684908,0.00076191686,0.0052907164,0.022461955,0.0111307],"genre_scores_gemma":[0.38006628,0.0029529843,0.57076937,0.000867302,0.00047485225,0.0008188623,0.01655329,0.0015786618,0.025918413],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989945,0.0002337404,0.000075551725,0.00030306118,0.00028697556,0.00010615484],"domain_scores_gemma":[0.9984309,0.0005557997,0.00009465473,0.00047729,0.00036815132,0.00007317225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00073288806,0.0015293661,0.0010747127,0.0023452612,0.00032377863,0.0013359523,0.0010167786,0.00095287105,0.005918308],"category_scores_gemma":[0.0040636174,0.00028002134,0.0009206209,0.0015127474,0.0006281721,0.004062784,0.002647349,0.0009019356,0.0037396885],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078637974,0.00035272198,0.00084413873,0.00094715937,0.00013031498,0.0002290596,0.00041677774,0.010636614,0.09423765,0.0045108967,0.033053923,0.8538544],"study_design_scores_gemma":[0.0002604748,0.0014228794,0.0072661582,0.0002477846,0.00040981328,0.0017159813,0.0013441425,0.6227817,0.23213725,0.036803067,0.09525526,0.00035549907],"about_ca_topic_score_codex":0.0027100258,"about_ca_topic_score_gemma":0.0038107142,"teacher_disagreement_score":0.005918308,"about_ca_system_score_codex":0.0004668874,"about_ca_system_score_gemma":0.00056186016,"threshold_uncertainty_score":0.019798696},"labels":[],"label_agreement":null},{"id":"W4404782772","doi":"10.18653/v1/2024.emnlp-main.723","title":"Story Morals: Surfacing value-driven narrative schemas using large language models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Narrative; Value (mathematics); Computer science; Linguistics; Natural language processing; Philosophy; Machine learning","score_opus":0.05163279700465704,"score_gpt":0.30954038913146326,"score_spread":0.2579075921268062,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782772","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.079897225,0.0008616359,0.89084,0.0013722255,0.00012354246,0.0005900472,0.010868731,0.010690537,0.0047560576],"genre_scores_gemma":[0.34847587,0.00034282752,0.6271201,0.00022931396,0.000056026758,0.00056077784,0.02068699,0.00077999313,0.0017481485],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976839,0.0013841464,0.00016299007,0.000476161,0.0002312245,0.00006168843],"domain_scores_gemma":[0.98919755,0.0078498535,0.00080686866,0.0012484132,0.0006963162,0.00020101988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034317626,0.0009868015,0.00040222536,0.0029472008,0.00075181597,0.003268827,0.0016168639,0.0014187531,0.0032270795],"category_scores_gemma":[0.021513015,0.0005686374,0.0014656996,0.001556688,0.00079631625,0.005794783,0.0021687525,0.002147124,0.0015521907],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007072527,0.00047818478,0.04680817,0.0026758378,0.0004803018,0.0011635105,0.025167773,0.12294656,0.019319942,0.09141609,0.051516753,0.6373196],"study_design_scores_gemma":[0.000049660084,0.00007182786,0.0043951734,0.00026360078,0.00006945237,0.0003243044,0.0030609006,0.8732757,0.006686629,0.069040425,0.042695187,0.0000670888],"about_ca_topic_score_codex":0.006508631,"about_ca_topic_score_gemma":0.012915566,"teacher_disagreement_score":0.006508631,"about_ca_system_score_codex":0.001346967,"about_ca_system_score_gemma":0.0010656132,"threshold_uncertainty_score":0.018149137},"labels":[],"label_agreement":null},{"id":"W4404782839","doi":"10.18653/v1/2024.emnlp-main.243","title":"STOP! Benchmarking Large Language Models with Sensitivity Testing on Offensive Progressions","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Offensive; Benchmarking; Computer science; Sensitivity (control systems); Reliability engineering; Engineering; Operations research; Electronic engineering; Business","score_opus":0.032574780700197675,"score_gpt":0.27298695207963375,"score_spread":0.24041217137943607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782839","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5994344,0.007263386,0.23091488,0.0044913623,0.0023709645,0.0012569487,0.04180084,0.09243698,0.020030195],"genre_scores_gemma":[0.7989014,0.0007313564,0.107402675,0.0018012034,0.0003731785,0.00075626874,0.08236494,0.0041021067,0.003567004],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98868805,0.006509674,0.00079691625,0.0021315461,0.0013549953,0.0005186884],"domain_scores_gemma":[0.9617374,0.028746849,0.0007600089,0.0051027676,0.002866745,0.0007862863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012517678,0.0034804146,0.0013345629,0.0024108798,0.0012154244,0.003105555,0.0030992418,0.0028740335,0.0062350016],"category_scores_gemma":[0.060581934,0.00081838144,0.0021894178,0.0017807676,0.0013904255,0.0056489646,0.0034548151,0.004873519,0.0047365734],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003055806,0.0018522384,0.049485426,0.0023732798,0.0021159046,0.0010912253,0.0018948235,0.43704498,0.01163797,0.009275074,0.17449322,0.30568007],"study_design_scores_gemma":[0.0003499083,0.00044027725,0.004217658,0.00017346068,0.00016380724,0.00028420473,0.0005142394,0.95824623,0.0074194972,0.013796105,0.014267222,0.00012741855],"about_ca_topic_score_codex":0.015838232,"about_ca_topic_score_gemma":0.022213152,"teacher_disagreement_score":0.015838232,"about_ca_system_score_codex":0.0015188169,"about_ca_system_score_gemma":0.0022544158,"threshold_uncertainty_score":0.066200614},"labels":[],"label_agreement":null},{"id":"W4404782882","doi":"10.18653/v1/2024.emnlp-main.764","title":"A Systematic Survey and Critical Review on Evaluating Large Language Models: Challenges, Limitations, and Recommendations","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Computer science; Data science; Management science; Engineering","score_opus":0.3056448804553619,"score_gpt":0.40923776009083845,"score_spread":0.10359287963547653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782882","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011171952,0.98528105,0.002291583,0.007995433,0.0008642494,0.0008104494,0.00082114176,0.000063585525,0.0007552786],"genre_scores_gemma":[0.021253396,0.95343393,0.010922116,0.008512832,0.0010790778,0.0032674002,0.0011890773,0.000100430865,0.00024169529],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.86232644,0.07470629,0.040324982,0.004779056,0.016889613,0.0009735853],"domain_scores_gemma":[0.3180421,0.5748593,0.03696271,0.011515818,0.056127753,0.0024923496],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.21987496,0.002579168,0.009298088,0.02384304,0.0016729349,0.0075739934,0.005649747,0.0028036109,0.00564239],"category_scores_gemma":[0.5153186,0.0023744735,0.011287931,0.015066783,0.0037008333,0.010412224,0.005156078,0.0040396503,0.0011555699],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079437724,0.00009936394,0.0034093694,0.6172892,0.015712945,0.00011661655,0.0013525364,0.000610511,0.000349457,0.0023201297,0.026004635,0.3319409],"study_design_scores_gemma":[0.00044534684,0.0005008483,0.0034362809,0.86482865,0.038290516,0.0002058518,0.0011662504,0.00049228646,0.00042603843,0.0042962264,0.08578683,0.00012488187],"about_ca_topic_score_codex":0.009251032,"about_ca_topic_score_gemma":0.027538342,"teacher_disagreement_score":0.21987496,"about_ca_system_score_codex":0.010436238,"about_ca_system_score_gemma":0.030105753,"threshold_uncertainty_score":0.96203303},"labels":[],"label_agreement":null},{"id":"W4404782892","doi":"10.18653/v1/2024.emnlp-main.250","title":"PromptReps: Prompting Large Language Models to Generate Dense and Sparse Representations for Zero-Shot Document Retrieval","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; University of Waterloo","keywords":"Zero (linguistics); Computer science; Shot (pellet); Natural language processing; Artificial intelligence; Language model; Information retrieval; Linguistics; Materials science","score_opus":0.06476820452331428,"score_gpt":0.33275380316027114,"score_spread":0.26798559863695687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404782892","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018788032,0.0004940772,0.9639849,0.00021368016,0.000115936484,0.00016794035,0.00048150032,0.014488469,0.0012655152],"genre_scores_gemma":[0.29690504,0.0005335297,0.6839032,0.00066309015,0.0002587124,0.00064054463,0.0047970302,0.0012410362,0.0110577345],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999047,0.0003239305,0.000060681185,0.00024556552,0.0002312456,0.00009151067],"domain_scores_gemma":[0.9982698,0.0007022652,0.00012414357,0.000522085,0.00028542633,0.00009633419],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001416388,0.0014640243,0.0012210633,0.0010971319,0.00045348864,0.0011420458,0.0023063833,0.0013032791,0.004308442],"category_scores_gemma":[0.0056001125,0.0005506129,0.0010871228,0.0009707942,0.00074139074,0.0038315041,0.002312667,0.0019678036,0.0039524683],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006339011,0.00048708933,0.0012187777,0.0005707678,0.00018682025,0.00028608053,0.0003850613,0.08603788,0.04548468,0.0141442865,0.03530804,0.8152566],"study_design_scores_gemma":[0.00008684145,0.00022126139,0.00030068084,0.000017256778,0.000033920925,0.00015881604,0.00008739586,0.9634942,0.018041512,0.011480314,0.0060279397,0.000049778664],"about_ca_topic_score_codex":0.004099715,"about_ca_topic_score_gemma":0.008311663,"teacher_disagreement_score":0.004308442,"about_ca_system_score_codex":0.0007937212,"about_ca_system_score_gemma":0.0012690205,"threshold_uncertainty_score":0.014413118},"labels":[],"label_agreement":null},{"id":"W4404783494","doi":"10.18653/v1/2024.emnlp-industry.10","title":"DL-QAT: Weight-Decomposed Low-Rank Quantization-Aware Training for Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Advanced Micro Devices (Canada)","funders":"","keywords":"Computer science; Quantization (signal processing); Rank (graph theory); Language model; Artificial intelligence; Natural language processing; Speech recognition; Mathematics; Algorithm; Combinatorics","score_opus":0.04104453252529369,"score_gpt":0.2967512469939502,"score_spread":0.25570671446865656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404783494","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01016297,0.0010592971,0.97442573,0.000542941,0.00022169988,0.0001182052,0.00050739426,0.011515029,0.0014467371],"genre_scores_gemma":[0.28391606,0.0007259638,0.6999878,0.0013681483,0.0003377846,0.00043897724,0.004819691,0.0016517339,0.0067537623],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981964,0.00067672064,0.00013193858,0.0004214236,0.00040785814,0.0001655676],"domain_scores_gemma":[0.99622285,0.0020630178,0.00016521425,0.00075490965,0.0006399732,0.00015401014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022618135,0.0019778542,0.0015646956,0.000908668,0.00078917574,0.0013985324,0.0026724061,0.0017267484,0.007211184],"category_scores_gemma":[0.013423531,0.0008072223,0.0011861274,0.0011943085,0.000910436,0.0034558973,0.0019329772,0.004364688,0.0039950116],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000538329,0.0003532125,0.0018798007,0.0004436501,0.00022041397,0.00024576107,0.0002898304,0.2613691,0.016412131,0.012120498,0.040197104,0.66593015],"study_design_scores_gemma":[0.00003492139,0.000057391662,0.00014068732,0.000014191134,0.000015161475,0.000045402317,0.000027974065,0.98759896,0.0030107885,0.00736814,0.001671615,0.00001481574],"about_ca_topic_score_codex":0.01180017,"about_ca_topic_score_gemma":0.023894908,"teacher_disagreement_score":0.01180017,"about_ca_system_score_codex":0.00096627587,"about_ca_system_score_gemma":0.0020620183,"threshold_uncertainty_score":0.024123788},"labels":[],"label_agreement":null},{"id":"W4404783705","doi":"10.18653/v1/2024.emnlp-main.211","title":"Decompose and Compare Consistency: Measuring VLMs’ Answer Reliability via Task-Decomposition Consistency Comparison","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Consistency (knowledge bases); Reliability (semiconductor); Computer science; Task (project management); Decomposition; Reliability engineering; Data mining; Artificial intelligence; Engineering","score_opus":0.031966096383067244,"score_gpt":0.2864189849729321,"score_spread":0.25445288858986487,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404783705","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48351246,0.001578962,0.48816222,0.00074778637,0.00022949377,0.0022667313,0.0036264297,0.008680138,0.011195796],"genre_scores_gemma":[0.82318336,0.00013853716,0.1686363,0.00034439916,0.00012623183,0.0017438933,0.0035635435,0.0008646897,0.0013990625],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9807798,0.009038879,0.0018783099,0.0031941894,0.0045876205,0.0005211378],"domain_scores_gemma":[0.84115565,0.11111152,0.016147252,0.013060939,0.016551515,0.001973092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022665523,0.0013454496,0.0010158428,0.0038571935,0.0006367365,0.0027961046,0.0012423275,0.0015105285,0.0034633183],"category_scores_gemma":[0.1950502,0.0005096719,0.0010743482,0.0018794977,0.00082996086,0.0036888164,0.0037768628,0.0016804049,0.0015089385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0052226684,0.0011521834,0.25235587,0.0020472745,0.0011442364,0.00022310464,0.00787474,0.01909174,0.032294884,0.00957304,0.014985438,0.65403485],"study_design_scores_gemma":[0.0011532826,0.0033161447,0.40044275,0.00061684696,0.0009970483,0.0009729548,0.0037349008,0.44091198,0.055753447,0.07055472,0.020849558,0.00069637765],"about_ca_topic_score_codex":0.0015258698,"about_ca_topic_score_gemma":0.0015593909,"teacher_disagreement_score":0.022665523,"about_ca_system_score_codex":0.00080121314,"about_ca_system_score_gemma":0.0012579958,"threshold_uncertainty_score":0.11986816},"labels":[],"label_agreement":null},{"id":"W4404783714","doi":"10.18653/v1/2024.emnlp-industry.86","title":"Query-OPT: Optimizing Inference of Large Language Models via Multi-Query Instructions in Meeting Summarization","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Stornoway Diamond (Canada)","funders":"","keywords":"Computer science; Automatic summarization; RDF query language; Query language; Inference; Query optimization; Web search query; Query expansion; Sargable; Information retrieval; Web query classification; Query by Example; Artificial intelligence; Search engine","score_opus":0.0252817177744218,"score_gpt":0.28866278496418896,"score_spread":0.26338106718976717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404783714","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05486945,0.0008716522,0.7923715,0.0009108652,0.00017834664,0.00035873146,0.0034565239,0.14502943,0.0019535872],"genre_scores_gemma":[0.35190567,0.00024112404,0.6279935,0.0005641897,0.00013872184,0.00033113844,0.011118813,0.0046579,0.0030488162],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974093,0.0011575418,0.0001633178,0.00068950217,0.00036800705,0.0002123328],"domain_scores_gemma":[0.99485517,0.0036018,0.00020040544,0.00071054784,0.00045370922,0.00017823988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003212179,0.0019800174,0.0012948354,0.0008434067,0.00068658945,0.001701648,0.0026006238,0.0014814199,0.0066514197],"category_scores_gemma":[0.013188602,0.00069284835,0.0012515313,0.0010893019,0.00062176096,0.0036094787,0.0019048153,0.002267607,0.0029786613],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028859605,0.0007206532,0.008811925,0.0013876403,0.00045011306,0.0005030476,0.0018829713,0.18936753,0.05767834,0.009903153,0.06494578,0.6614629],"study_design_scores_gemma":[0.00015818472,0.00018960617,0.00088569697,0.000014212464,0.000079313555,0.000080327576,0.0003864896,0.9687066,0.016597992,0.006102631,0.0067570587,0.00004181983],"about_ca_topic_score_codex":0.013891774,"about_ca_topic_score_gemma":0.02381394,"teacher_disagreement_score":0.013891774,"about_ca_system_score_codex":0.0011163928,"about_ca_system_score_gemma":0.0025308356,"threshold_uncertainty_score":0.027621865},"labels":[],"label_agreement":null},{"id":"W4404785115","doi":"10.1007/s40593-024-00441-x","title":"Predicting Tags for Learner Questions on Stack Overflow","year":2024,"lang":"en","type":"article","venue":"International Journal of Artificial Intelligence in Education","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Killam Trusts","keywords":"Stack (abstract data type); Computer science; Educational technology; Multimedia; Mathematics education; Programming language; Mathematics","score_opus":0.05232476163451524,"score_gpt":0.38408274498883294,"score_spread":0.3317579833543177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404785115","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94733727,0.0010764531,0.04231724,0.0004875197,0.00019299815,0.0001949598,0.0026511296,0.0035739602,0.0021684803],"genre_scores_gemma":[0.9667476,0.00020785519,0.024989042,0.00012083687,0.00008368257,0.00006849358,0.0043005906,0.000089649184,0.0033924088],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991757,0.00028367632,0.00005991512,0.00019695029,0.00017318723,0.00011055266],"domain_scores_gemma":[0.9937065,0.003920263,0.0004982853,0.00027896996,0.0011526541,0.00044347285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016806385,0.001113916,0.00058242155,0.0035297845,0.00053303345,0.0010472713,0.00058194157,0.0017412961,0.002291244],"category_scores_gemma":[0.009417695,0.00020645674,0.00070605567,0.0009826861,0.00033878657,0.0024339953,0.001021787,0.0011729805,0.0017193904],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003776779,0.0015305205,0.2881165,0.0011258499,0.00025258347,0.0010685581,0.003292697,0.03386139,0.059333727,0.0025189104,0.027611222,0.5775112],"study_design_scores_gemma":[0.000071845934,0.0010428166,0.10983766,0.00016155948,0.00020233821,0.00032497474,0.0019317662,0.8275145,0.041752163,0.00521865,0.011824839,0.00011683552],"about_ca_topic_score_codex":0.006472692,"about_ca_topic_score_gemma":0.009460349,"teacher_disagreement_score":0.006472692,"about_ca_system_score_codex":0.0008326303,"about_ca_system_score_gemma":0.0007236293,"threshold_uncertainty_score":0.012870073},"labels":[],"label_agreement":null},{"id":"W4404787914","doi":"10.1109/access.2024.3507382","title":"Enhancing Sindhi Word Segmentation Using Subword Representation Learning and Position-Aware Self-Attention","year":2024,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"National Key Research and Development Program of China","keywords":"Computer science; Natural language processing; Representation (politics); Segmentation; Artificial intelligence; Speech recognition; Word (group theory); Text segmentation; Position (finance); Linguistics","score_opus":0.033268435060763706,"score_gpt":0.34569485629040925,"score_spread":0.31242642122964553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404787914","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30888492,0.0061890944,0.6229783,0.0007904312,0.0007294301,0.0004477199,0.003953099,0.04198951,0.014037451],"genre_scores_gemma":[0.60212916,0.0012798482,0.35574922,0.0006111951,0.00028268047,0.00040203254,0.018626185,0.0012590162,0.019660696],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944085,0.00009148652,0.000049039823,0.00026545717,0.0000837058,0.00006937231],"domain_scores_gemma":[0.9991899,0.00028379206,0.0000851801,0.00018951707,0.00020725105,0.00004441275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006646829,0.0017558034,0.0012694164,0.0022387316,0.0006166644,0.0011252301,0.0013407021,0.001293386,0.003145849],"category_scores_gemma":[0.0017898319,0.00038369064,0.001016535,0.0022128767,0.00045583356,0.0032110172,0.0012405851,0.0014435931,0.004377891],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004886135,0.00033723225,0.00289583,0.0004962319,0.00014987412,0.00022533514,0.00042377453,0.020049538,0.0556511,0.0026561415,0.022113388,0.89451283],"study_design_scores_gemma":[0.00006494264,0.00029286277,0.004132352,0.000043163906,0.00014869387,0.00031260643,0.00033865264,0.9216059,0.053807463,0.006693912,0.01250083,0.00005865383],"about_ca_topic_score_codex":0.0065894877,"about_ca_topic_score_gemma":0.014287298,"teacher_disagreement_score":0.0065894877,"about_ca_system_score_codex":0.0007332014,"about_ca_system_score_gemma":0.0011507082,"threshold_uncertainty_score":0.013102293},"labels":[],"label_agreement":null},{"id":"W4404788057","doi":"10.1145/3652892.3700758","title":"Menos: Split Fine-Tuning Large Language Models with Efficient GPU Memory Sharing","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Parallel computing; Fine-tuning; Computational science; Computer graphics (images); Physics","score_opus":0.02032411432330952,"score_gpt":0.2499054059051806,"score_spread":0.22958129158187107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404788057","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.067877054,0.0009302784,0.8641188,0.00040695036,0.00028945523,0.00021280449,0.00053673773,0.059037182,0.0065906853],"genre_scores_gemma":[0.5910603,0.0002721546,0.39400783,0.0007206319,0.00010646167,0.0004175356,0.002122003,0.0041701016,0.0071230023],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989631,0.00021255917,0.00006130035,0.00032496272,0.00028050772,0.00015746168],"domain_scores_gemma":[0.998741,0.00039452885,0.00005815545,0.0005228852,0.00017945902,0.00010408899],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012570011,0.0015136899,0.0010115731,0.00054262,0.00061290327,0.0015935305,0.0034594561,0.0010521404,0.005130267],"category_scores_gemma":[0.0055184443,0.00073743064,0.0011843921,0.0006042606,0.00084303226,0.0028200506,0.0029389146,0.0025442578,0.0024026427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010986181,0.0006477516,0.00636456,0.00026218445,0.00034545464,0.0003664725,0.0005920167,0.4940175,0.030968135,0.014867664,0.039457917,0.4110117],"study_design_scores_gemma":[0.00005980826,0.000061231534,0.00022486984,0.000006448496,0.000017009643,0.000036454214,0.00004405575,0.98413545,0.005667784,0.006050801,0.0036808273,0.000015253483],"about_ca_topic_score_codex":0.009736564,"about_ca_topic_score_gemma":0.018192898,"teacher_disagreement_score":0.009736564,"about_ca_system_score_codex":0.0011318199,"about_ca_system_score_gemma":0.0021887962,"threshold_uncertainty_score":0.019359767},"labels":[],"label_agreement":null},{"id":"W4404851119","doi":"10.1016/j.eswa.2024.125924","title":"Pairwise dual-level alignment for cross-prompt automated essay scoring","year":2024,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada Research Chairs; University of Toronto","funders":"Taishan Scholar Project of Shandong Province; Natural Science Foundation of Shandong Province; National Natural Science Foundation of China","keywords":"Pairwise comparison; Dual (grammatical number); Computer science; Artificial intelligence; Cross-validation; Data mining; Natural language processing; Pattern recognition (psychology)","score_opus":0.04273792357509929,"score_gpt":0.3207816349410548,"score_spread":0.2780437113659555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404851119","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024395388,0.0005568009,0.94456565,0.00020943314,0.00040970228,0.00028985547,0.002119701,0.02171845,0.0057349806],"genre_scores_gemma":[0.26531228,0.00020718151,0.7112133,0.00017488876,0.00022470237,0.0006217745,0.010946666,0.0030066767,0.0082924785],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99220383,0.0031421562,0.00063602603,0.0020722959,0.001310861,0.00063482346],"domain_scores_gemma":[0.9867264,0.004715504,0.0007212324,0.0023441787,0.0047997725,0.00069286523],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045725508,0.0015663204,0.0016190321,0.003968651,0.0020371461,0.0032883768,0.0024490603,0.0021743483,0.0171283],"category_scores_gemma":[0.022230895,0.00093558105,0.0010025029,0.0038833346,0.0005705011,0.003068874,0.0050858147,0.0032689294,0.016809154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000962982,0.00038446797,0.003875869,0.00055669324,0.0001575169,0.00023939474,0.0009241128,0.0062953876,0.05327974,0.0077091707,0.04031002,0.8853047],"study_design_scores_gemma":[0.00028411055,0.0007982772,0.010983554,0.00024788314,0.00029602443,0.0010251828,0.0019762942,0.73788714,0.108667836,0.053125534,0.08444628,0.00026184827],"about_ca_topic_score_codex":0.0017349215,"about_ca_topic_score_gemma":0.0045644254,"teacher_disagreement_score":0.0171283,"about_ca_system_score_codex":0.00073869,"about_ca_system_score_gemma":0.002787334,"threshold_uncertainty_score":0.057299852},"labels":[],"label_agreement":null},{"id":"W4404875179","doi":"10.1016/j.engappai.2024.109490","title":"Low-cost language models: Survey and performance evaluation on Python code generation","year":2024,"lang":"en","type":"article","venue":"Engineering Applications of Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Novelis (Canada)","funders":"","keywords":"Computer science; Python (programming language); Programming language; Code generation; Operating system","score_opus":0.08149017670230732,"score_gpt":0.3154793380580967,"score_spread":0.23398916135578937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404875179","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.081403956,0.008260959,0.47296384,0.0017434645,0.0008683954,0.0012132741,0.00892823,0.4055501,0.019067738],"genre_scores_gemma":[0.31749895,0.00719536,0.57618886,0.0014768934,0.00018736292,0.0010625076,0.039519176,0.047243353,0.0096275285],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957371,0.0010644029,0.00043184563,0.0005308118,0.0018747636,0.0003610472],"domain_scores_gemma":[0.9887721,0.0050248927,0.000559446,0.0032215703,0.0020237335,0.0003982521],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033469542,0.0018976253,0.0010236021,0.0024126044,0.0008036853,0.0023861956,0.007176474,0.0010646649,0.009950224],"category_scores_gemma":[0.02136119,0.001372321,0.0020241386,0.0037605444,0.000954703,0.005083388,0.0025853922,0.0026596969,0.0072819768],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015351993,0.0009962524,0.009366813,0.0042398083,0.00048397022,0.0003087229,0.0006250876,0.06688216,0.014186506,0.018706834,0.15690246,0.7257661],"study_design_scores_gemma":[0.0008292026,0.0008169082,0.0040903217,0.00074128516,0.0005289641,0.000744704,0.0003686918,0.7636869,0.059280183,0.023070391,0.14559297,0.00024950693],"about_ca_topic_score_codex":0.014816572,"about_ca_topic_score_gemma":0.014352651,"teacher_disagreement_score":0.014816572,"about_ca_system_score_codex":0.0018331932,"about_ca_system_score_gemma":0.0049260105,"threshold_uncertainty_score":0.03328681},"labels":[],"label_agreement":null},{"id":"W4404917089","doi":"10.2196/60334","title":"Chinese Clinical Named Entity Recognition With Segmentation Synonym Sentence Synthesis Mechanism: Algorithm Development and Validation","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Conditional random field; Natural language processing; Named-entity recognition; Sentence; Vocabulary; Segmentation; Synonym (taxonomy); Machine learning; Task (project management)","score_opus":0.030130818809666185,"score_gpt":0.3196210585014545,"score_spread":0.28949023969178833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404917089","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25579715,0.0026257539,0.691347,0.0011750156,0.0004899476,0.0021841757,0.0052825985,0.036288776,0.004809613],"genre_scores_gemma":[0.3737771,0.00071933115,0.60113585,0.00041120482,0.00009285602,0.0017156198,0.01807719,0.00044354633,0.003627308],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984133,0.0004011846,0.00019750434,0.00060689764,0.00027196525,0.00010921152],"domain_scores_gemma":[0.99685645,0.0014538122,0.00016885574,0.00047418452,0.00093905645,0.00010758766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041017714,0.001294812,0.001138065,0.00149531,0.0007626925,0.001002864,0.0025816038,0.001508713,0.004684125],"category_scores_gemma":[0.007626805,0.00040747598,0.00097225397,0.0012482337,0.00052127347,0.0019911986,0.0013690536,0.0016143795,0.00227706],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008692971,0.0006150153,0.008509499,0.0005965041,0.0003260368,0.0003770989,0.00023736834,0.13456038,0.0124524925,0.0028743125,0.022552654,0.81602937],"study_design_scores_gemma":[0.00014132306,0.00018048917,0.0017575709,0.000026870002,0.000059157235,0.00015326466,0.0001023026,0.98069936,0.012365071,0.0013187018,0.0031681233,0.000027688257],"about_ca_topic_score_codex":0.016792346,"about_ca_topic_score_gemma":0.012752131,"teacher_disagreement_score":0.016792346,"about_ca_system_score_codex":0.0013488007,"about_ca_system_score_gemma":0.0034005546,"threshold_uncertainty_score":0.03338921},"labels":[],"label_agreement":null},{"id":"W4404932874","doi":"10.1007/978-981-96-0567-5_15","title":"Vector Representation Learning of Skills for Collaborative Team Recommendation: A Comparative Study","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Computer science; Representation (politics); Artificial intelligence; Human–computer interaction; Knowledge management","score_opus":0.04193913212345869,"score_gpt":0.3319939661844991,"score_spread":0.29005483406104043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404932874","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.728847,0.002689155,0.25386772,0.00078380137,0.00015745014,0.00031694662,0.00043691858,0.0015193278,0.011381687],"genre_scores_gemma":[0.95069665,0.0006521276,0.04424835,0.000056202356,0.00004492871,0.000107885426,0.0006408257,0.000081137034,0.00347187],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981469,0.0009768298,0.00008928491,0.0003240599,0.00033404006,0.00012901466],"domain_scores_gemma":[0.96898234,0.026825385,0.00059238484,0.0016286878,0.0015080278,0.0004632195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039865165,0.000524899,0.00079331925,0.0012778109,0.00039040507,0.001754373,0.0013664647,0.0011700679,0.003475313],"category_scores_gemma":[0.022262385,0.00020365987,0.0007119848,0.0017839117,0.0003504386,0.0040231314,0.000849113,0.0014396714,0.00091733615],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015542333,0.0023529492,0.013779833,0.00033460316,0.00023230816,0.000028050468,0.000660585,0.04457355,0.0020298415,0.0033176648,0.0035340784,0.9276023],"study_design_scores_gemma":[0.00017583572,0.002067985,0.016445093,0.00008434695,0.00028758717,0.000085837964,0.0008942927,0.963835,0.0034848629,0.009272456,0.0033103747,0.00005629488],"about_ca_topic_score_codex":0.008025565,"about_ca_topic_score_gemma":0.0047337287,"teacher_disagreement_score":0.008025565,"about_ca_system_score_codex":0.0011728235,"about_ca_system_score_gemma":0.0010411588,"threshold_uncertainty_score":0.021082997},"labels":[],"label_agreement":null},{"id":"W4404960594","doi":"10.3390/en17236063","title":"Domain-Specific Large Language Model for Renewable Energy and Hydrogen Deployment Strategies","year":2024,"lang":"en","type":"article","venue":"Energies","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Renewable energy; Software deployment; Computer science; Domain (mathematical analysis); Environmental economics; Work (physics); Engineering; Economics; Mechanical engineering; Electrical engineering","score_opus":0.017332663688469176,"score_gpt":0.2462847393560578,"score_spread":0.22895207566758863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404960594","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1325761,0.0026034287,0.8286638,0.004674029,0.0004979807,0.0003342205,0.009482536,0.01016308,0.011004778],"genre_scores_gemma":[0.79014575,0.00071586075,0.18493406,0.00088963297,0.00022099506,0.0006904449,0.0130647225,0.0007824218,0.008555954],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993325,0.00034252342,0.000036437163,0.00019391249,0.000048741716,0.000045842477],"domain_scores_gemma":[0.99676114,0.0026571397,0.00011547631,0.00015337304,0.00022236937,0.00009052082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016850622,0.001025506,0.00048605938,0.0010954974,0.00044742163,0.0013200645,0.0013130183,0.0015141496,0.0050444305],"category_scores_gemma":[0.0062716347,0.0004626305,0.0010980185,0.00077826506,0.00041315774,0.0022768886,0.0010491128,0.002769202,0.0028329082],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005030277,0.00034163875,0.009599844,0.0005828091,0.00022881669,0.0005038177,0.0009677244,0.7113943,0.0073606204,0.024699293,0.038399283,0.20541884],"study_design_scores_gemma":[0.000016901255,0.0000200374,0.00042955342,0.000020688407,0.000018369386,0.000036819154,0.00006049098,0.9855592,0.0007325111,0.009577761,0.0035158538,0.000011801261],"about_ca_topic_score_codex":0.0072243083,"about_ca_topic_score_gemma":0.01661039,"teacher_disagreement_score":0.0072243083,"about_ca_system_score_codex":0.0012140885,"about_ca_system_score_gemma":0.0011431719,"threshold_uncertainty_score":0.016875327},"labels":[],"label_agreement":null},{"id":"W4404971593","doi":"10.1007/978-3-031-78498-9_3","title":"ConCSE: Unified Contrastive Learning and Augmentation for Code-Switched Embeddings","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Code (set theory); Programming language; Artificial intelligence; Natural language processing","score_opus":0.020165552351568583,"score_gpt":0.2800862608764734,"score_spread":0.25992070852490484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404971593","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0075563584,0.00050570664,0.97829974,0.00020887076,0.00026518505,0.000109504595,0.0004888813,0.008750153,0.003815629],"genre_scores_gemma":[0.1828413,0.0005514224,0.79033065,0.00055556,0.0002935835,0.00046200468,0.004350928,0.002529961,0.01808466],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99878186,0.0002979304,0.0000524233,0.00042253,0.00032271398,0.00012260645],"domain_scores_gemma":[0.9981698,0.00070104166,0.00006666615,0.0006188376,0.0003444879,0.000099250086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014457963,0.0020379557,0.0015503361,0.001243769,0.0007856042,0.001996115,0.0034911812,0.0021622716,0.011411284],"category_scores_gemma":[0.005938935,0.00076795917,0.0014733487,0.0013121507,0.0012793976,0.0049240314,0.0051819123,0.0042237407,0.005734984],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005054309,0.0003120588,0.000624015,0.0003243914,0.000102285565,0.00013851805,0.00015066596,0.06135852,0.016370302,0.051807743,0.034023024,0.834283],"study_design_scores_gemma":[0.000045395773,0.00012261653,0.00015907075,0.000038273476,0.000025034376,0.00009003592,0.00004124913,0.9233314,0.010396709,0.055571757,0.010149861,0.0000285372],"about_ca_topic_score_codex":0.0033448203,"about_ca_topic_score_gemma":0.00814426,"teacher_disagreement_score":0.011411284,"about_ca_system_score_codex":0.00088553195,"about_ca_system_score_gemma":0.0016187417,"threshold_uncertainty_score":0.03817451},"labels":[],"label_agreement":null},{"id":"W4405127663","doi":"10.1007/s41060-024-00693-9","title":"AI-generated or AI touch-up? Identifying AI contribution in text data","year":2024,"lang":"en","type":"article","venue":"International Journal of Data Science and Analytics","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Data science","score_opus":0.1328814845314103,"score_gpt":0.4163926192477443,"score_spread":0.283511134716334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405127663","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6913803,0.010885006,0.21234119,0.015301673,0.001921841,0.00063550164,0.004330489,0.002930136,0.06027382],"genre_scores_gemma":[0.9549793,0.0011655504,0.03401045,0.0008398809,0.0008884791,0.00027402272,0.0028542308,0.00057648605,0.00441166],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9839977,0.008135396,0.0010670517,0.0021978954,0.003917029,0.0006848919],"domain_scores_gemma":[0.8047562,0.1559989,0.008271168,0.01199719,0.015291948,0.003684607],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014001915,0.00079272,0.00093608437,0.010014469,0.0021422969,0.008445439,0.0019136346,0.0023043281,0.004941438],"category_scores_gemma":[0.16804515,0.0006056066,0.0007207776,0.010700061,0.0026818032,0.014857933,0.006041397,0.0030721945,0.0020125294],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017914267,0.0006115535,0.25813323,0.0026116138,0.00075252895,0.0018062607,0.044722524,0.0061264164,0.012254319,0.09306535,0.030910306,0.5472145],"study_design_scores_gemma":[0.0002308968,0.00053737056,0.16105756,0.0019099639,0.0012865567,0.003022018,0.043260157,0.23988128,0.017628249,0.3585889,0.17230494,0.00029213377],"about_ca_topic_score_codex":0.0025254064,"about_ca_topic_score_gemma":0.0030100057,"teacher_disagreement_score":0.9859981,"about_ca_system_score_codex":0.0014758321,"about_ca_system_score_gemma":0.0019857914,"threshold_uncertainty_score":0.07405013},"labels":[],"label_agreement":null},{"id":"W4405144017","doi":"10.1145/3673791.3698440","title":"Paradigm Shifts in Team Recommendation: From Historical Subgraph Optimization to Emerging Graph Neural Network","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"Universitas Brawijaya","keywords":"Computer science; Induced subgraph isomorphism problem; Subgraph isomorphism problem; Artificial neural network; Graph; Artificial intelligence; Theoretical computer science; Data science; Line graph; Voltage graph","score_opus":0.023698197747669285,"score_gpt":0.25338985317252954,"score_spread":0.22969165542486025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405144017","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022156805,0.005652487,0.96221226,0.0021065557,0.00024081916,0.000071859096,0.00044050446,0.0007253055,0.0063934387],"genre_scores_gemma":[0.43987477,0.008153588,0.5387055,0.0012839172,0.000794658,0.00026983672,0.0019761615,0.0006621708,0.008279413],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99903595,0.00042108374,0.000034161614,0.00028560753,0.00015702535,0.000066247194],"domain_scores_gemma":[0.99770844,0.0013887197,0.00016130714,0.00036829917,0.00024033982,0.00013286795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018327178,0.0008641639,0.00095814123,0.0011732498,0.0006355952,0.0017193003,0.0021027185,0.0014964202,0.0025993774],"category_scores_gemma":[0.0077615357,0.00058919477,0.0009797375,0.0020125916,0.0009283623,0.0035603803,0.0013667099,0.0027739173,0.000829947],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021598028,0.0002396715,0.0043507605,0.0006684025,0.00028043083,0.000114750575,0.0004989617,0.46514633,0.0020254161,0.11655657,0.03083004,0.37907273],"study_design_scores_gemma":[0.000019153345,0.000040402112,0.0004994653,0.00006123666,0.000030507936,0.000044519107,0.00007969669,0.8835346,0.00043205172,0.108157225,0.0070852526,0.000015989795],"about_ca_topic_score_codex":0.009161762,"about_ca_topic_score_gemma":0.015428267,"teacher_disagreement_score":0.009161762,"about_ca_system_score_codex":0.0015095918,"about_ca_system_score_gemma":0.0010091129,"threshold_uncertainty_score":0.018216908},"labels":[],"label_agreement":null},{"id":"W4405183082","doi":"10.1145/3658644.3670334","title":"zkLLM: Zero Knowledge Proofs for Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Zero-knowledge proof; Computer science; Mathematical proof; Zero (linguistics); Programming language; Mathematics; Linguistics; Algorithm; Philosophy; Cryptography","score_opus":0.03503480681345038,"score_gpt":0.3043688524513657,"score_spread":0.2693340456379153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405183082","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003852373,0.0017414396,0.9616649,0.0046960125,0.00042290194,0.0002501212,0.0020633172,0.010529206,0.014779723],"genre_scores_gemma":[0.40333977,0.003732948,0.546391,0.005149177,0.0015610346,0.0010665039,0.007610716,0.005116217,0.026032763],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9880128,0.0054062936,0.00071901997,0.0013657582,0.0037604335,0.0007356146],"domain_scores_gemma":[0.9519571,0.03636817,0.001916035,0.007126239,0.0020147979,0.0006176613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010083511,0.0014005536,0.0015355167,0.0034139403,0.0022056387,0.008017889,0.0046924464,0.002919454,0.030250417],"category_scores_gemma":[0.06024309,0.0015793759,0.0025708266,0.0028827377,0.003708472,0.016734382,0.010730894,0.006831981,0.012030554],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025906402,0.000103500126,0.00043077886,0.0009278846,0.00013953255,0.0004351791,0.0004642486,0.0142266145,0.0017264776,0.8197771,0.055769715,0.105739936],"study_design_scores_gemma":[0.00007514815,0.000025479267,0.00006797037,0.000115466806,0.00003566027,0.00020747169,0.00004887298,0.056262046,0.0018977494,0.91610295,0.025127798,0.00003338144],"about_ca_topic_score_codex":0.0014097334,"about_ca_topic_score_gemma":0.002487714,"teacher_disagreement_score":0.030250417,"about_ca_system_score_codex":0.0031110486,"about_ca_system_score_gemma":0.0048619905,"threshold_uncertainty_score":0.10119778},"labels":[],"label_agreement":null},{"id":"W4405235374","doi":"10.1016/j.eswa.2024.126130","title":"Joint entity and relation extraction with table filling based on graph convolutional Networks","year":2024,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Novelis (Canada)","funders":"","keywords":"Computer science; Relationship extraction; Joint (building); Graph; Table (database); Relation (database); Artificial intelligence; Data mining; Pattern recognition (psychology); Theoretical computer science","score_opus":0.017283427904659514,"score_gpt":0.23851524714268935,"score_spread":0.22123181923802984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405235374","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043385725,0.0022885303,0.9003984,0.00068714307,0.00035273572,0.00038207718,0.015923135,0.030731525,0.0058508674],"genre_scores_gemma":[0.29882008,0.001518433,0.6410229,0.0002634335,0.00019225963,0.00027262454,0.04501647,0.0009913284,0.0119025335],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991202,0.00007690163,0.00007928413,0.0004226155,0.00018075635,0.00012014575],"domain_scores_gemma":[0.99883324,0.00048481615,0.000094110575,0.00029763533,0.00023035871,0.000059833783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006883092,0.0015699839,0.0013737071,0.0046546375,0.0008510193,0.0017827216,0.0017105354,0.0011746545,0.006957787],"category_scores_gemma":[0.002388368,0.00070529233,0.002051131,0.0055716815,0.0003924963,0.0040022316,0.0014940866,0.001518949,0.005003856],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005617625,0.000314474,0.0048249844,0.0005773803,0.00027181543,0.00047854567,0.00024942541,0.02130968,0.028850157,0.010833085,0.041784305,0.88994443],"study_design_scores_gemma":[0.00005593504,0.00015208841,0.005911811,0.0001262224,0.00045621363,0.0005186098,0.00025587535,0.85775965,0.050599985,0.044634014,0.039426308,0.00010330572],"about_ca_topic_score_codex":0.018953864,"about_ca_topic_score_gemma":0.037290417,"teacher_disagreement_score":0.018953864,"about_ca_system_score_codex":0.0010294149,"about_ca_system_score_gemma":0.0023854792,"threshold_uncertainty_score":0.037687063},"labels":[],"label_agreement":null},{"id":"W4405266245","doi":"10.31219/osf.io/nbm4f","title":"CoordiLang: Assessing Multi-Agent Coordination Skills in Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Université de Montréal","funders":"","keywords":"Benchmark (surveying); Computer science; Coordination game; Motor coordination; Comprehension; Inference; Cognitive science; Knowledge management; Artificial intelligence; Psychology; Microeconomics","score_opus":0.03513912502458995,"score_gpt":0.32375436779341704,"score_spread":0.2886152427688271,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405266245","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5431966,0.0011094665,0.41500005,0.0010792615,0.0002672167,0.0011032297,0.004479182,0.015527233,0.01823765],"genre_scores_gemma":[0.7997583,0.00015768799,0.19193895,0.0002075424,0.000025117124,0.00058769283,0.0053777806,0.00044197406,0.0015050102],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950348,0.0029055332,0.00034775527,0.0007868133,0.0007137285,0.00021140513],"domain_scores_gemma":[0.9775718,0.017129393,0.0012525739,0.0020117338,0.0010388762,0.0009955998],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006127919,0.0015799378,0.0006927989,0.0013960294,0.00063207094,0.002707537,0.0025588667,0.001980791,0.0035171274],"category_scores_gemma":[0.03351036,0.00044100717,0.0010839691,0.00074870046,0.0011405265,0.0035450328,0.003256213,0.0021928905,0.0009757464],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001475171,0.0017474896,0.025972243,0.0012993449,0.0005583233,0.00041781343,0.0015345239,0.7454161,0.011198966,0.019173736,0.013896678,0.17730962],"study_design_scores_gemma":[0.00010031448,0.00043306645,0.0021629476,0.000041697032,0.000039496936,0.000079466,0.00033691712,0.9805859,0.003429533,0.010036702,0.00270927,0.000044859047],"about_ca_topic_score_codex":0.010766462,"about_ca_topic_score_gemma":0.0119779585,"teacher_disagreement_score":0.010766462,"about_ca_system_score_codex":0.0014335962,"about_ca_system_score_gemma":0.0026224225,"threshold_uncertainty_score":0.03240794},"labels":[],"label_agreement":null},{"id":"W4405507526","doi":"10.21203/rs.3.rs-5453999/v1","title":"MIRACLE - Medical Information Retrieval using Clinical Language Embeddings for Retrieval Augmented Generation at the point of care","year":2024,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Miracle; Point (geometry); Information retrieval; Medical information; Point of care; Computer science; Natural language processing; Medicine; Mathematics; Political science; Nursing","score_opus":0.1254270944562968,"score_gpt":0.47308426116980823,"score_spread":0.34765716671351143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405507526","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031488612,0.0047092903,0.90597624,0.0034970385,0.0016484942,0.00075465144,0.011310735,0.031909365,0.008705627],"genre_scores_gemma":[0.31807277,0.001613066,0.60206914,0.0016966759,0.001416681,0.0009997003,0.04009326,0.002156016,0.031882666],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980672,0.0008403545,0.00010005391,0.0004810903,0.00032219136,0.0001890784],"domain_scores_gemma":[0.99749905,0.0011727758,0.000096986754,0.00058010756,0.0005176227,0.00013345973],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002658412,0.001462681,0.0014021398,0.001963095,0.0007191542,0.001980746,0.0016787268,0.0025399064,0.012481927],"category_scores_gemma":[0.008452442,0.00068850076,0.0019696644,0.0010840265,0.00050737336,0.0022477882,0.002745484,0.0020042907,0.010217273],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016898194,0.0007148604,0.0025837522,0.000766173,0.00041120316,0.00033145346,0.0002480549,0.051120915,0.017768398,0.0112433685,0.12570256,0.7874194],"study_design_scores_gemma":[0.0003431835,0.00060717156,0.0013232613,0.000096561285,0.00017357113,0.0005171228,0.00012127645,0.933908,0.01665267,0.019138766,0.027025297,0.000093171395],"about_ca_topic_score_codex":0.0047864392,"about_ca_topic_score_gemma":0.007369574,"teacher_disagreement_score":0.012481927,"about_ca_system_score_codex":0.00076126854,"about_ca_system_score_gemma":0.0021678363,"threshold_uncertainty_score":0.041756213},"labels":[],"label_agreement":null},{"id":"W4405597099","doi":"10.1038/s41746-024-01366-4","title":"Probabilistic medical predictions of large language models","year":2024,"lang":"en","type":"article","venue":"npj Digital Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"National Institutes of Health","keywords":"Probabilistic logic; Computer science; Artificial intelligence","score_opus":0.019218859412394984,"score_gpt":0.2821147625464776,"score_spread":0.2628959031340826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405597099","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.119364366,0.0030581607,0.8511013,0.006734556,0.00044125642,0.00051538437,0.007988123,0.0061572324,0.004639745],"genre_scores_gemma":[0.80236244,0.0012264107,0.18086967,0.0012732754,0.0005530397,0.00062184554,0.010142704,0.000617623,0.002333022],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99346197,0.0043315724,0.00037272304,0.0010184357,0.0005889118,0.00022634353],"domain_scores_gemma":[0.9189088,0.074575126,0.0019522812,0.0019509636,0.0020992244,0.00051357487],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012959622,0.0013554352,0.00082174956,0.0026867494,0.00048555623,0.0031305025,0.0015076215,0.0014524333,0.0040042964],"category_scores_gemma":[0.07083464,0.0005729426,0.0014793503,0.0014126766,0.0007465183,0.003433575,0.0019185115,0.0028300122,0.0021260483],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001673979,0.0003734933,0.049856313,0.0017496668,0.0006398141,0.00089215726,0.001750987,0.49150994,0.0042454037,0.0413312,0.041649885,0.36432725],"study_design_scores_gemma":[0.00008279405,0.00007221633,0.0020973743,0.00015120882,0.000080393904,0.00014920864,0.000113272305,0.9155532,0.0015594027,0.07574017,0.004357489,0.000043238917],"about_ca_topic_score_codex":0.0031313375,"about_ca_topic_score_gemma":0.0044938084,"teacher_disagreement_score":0.012959622,"about_ca_system_score_codex":0.0011830915,"about_ca_system_score_gemma":0.0015466117,"threshold_uncertainty_score":0.06853783},"labels":[],"label_agreement":null},{"id":"W4405618562","doi":"10.1111/emip.12663","title":"Instruction‐Tuned Large‐Language Models for Quality Control in Automatic Item Generation: A Feasibility Study","year":2024,"lang":"en","type":"article","venue":"Educational Measurement Issues and Practice","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Quality (philosophy); Computer science; Control (management); Item response theory; Language proficiency; Mathematics education; Natural language processing; Psychology; Artificial intelligence; Psychometrics; Developmental psychology","score_opus":0.19670056213224912,"score_gpt":0.4278503410420334,"score_spread":0.23114977890978428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405618562","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85047245,0.00013059871,0.14168076,0.00048335103,0.00006182345,0.0020635654,0.00029381318,0.0034825443,0.0013309424],"genre_scores_gemma":[0.86363566,0.00002986427,0.13459471,0.00011443701,0.000015455598,0.0008314723,0.00023366376,0.0002096155,0.0003351135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98491645,0.011454172,0.0009187471,0.0011373797,0.0012523144,0.00032097133],"domain_scores_gemma":[0.8163837,0.1545434,0.0035610131,0.012266138,0.011611646,0.0016340928],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034051638,0.0011578486,0.00075509393,0.000753626,0.000533775,0.0018844286,0.003070298,0.0013260817,0.0027079298],"category_scores_gemma":[0.13227774,0.000945216,0.0005778416,0.0007102212,0.000917786,0.00323792,0.0016269346,0.0021139183,0.0008768544],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011470159,0.020636007,0.086368434,0.0011612705,0.00043527715,0.0009060684,0.007445504,0.17333893,0.062251724,0.0055118543,0.0082976995,0.6221771],"study_design_scores_gemma":[0.0014510036,0.0047098524,0.0165485,0.00009154419,0.00015297103,0.00019765382,0.00053760735,0.94508886,0.026387118,0.0022199687,0.002486016,0.00012889982],"about_ca_topic_score_codex":0.008011228,"about_ca_topic_score_gemma":0.005956651,"teacher_disagreement_score":0.034051638,"about_ca_system_score_codex":0.001603922,"about_ca_system_score_gemma":0.0018809662,"threshold_uncertainty_score":0.18008447},"labels":[],"label_agreement":null},{"id":"W4405622254","doi":"10.32920/28072193","title":"Enhanced Biomedical Factoid Question Answering through Biomedical Knowledge Integration","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Question answering; Computer science; Information retrieval; Natural language processing","score_opus":0.034261900193245994,"score_gpt":0.32936014877687625,"score_spread":0.29509824858363026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405622254","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06754349,0.0040677995,0.883625,0.00338834,0.00036953835,0.0006883232,0.0063536735,0.026323212,0.0076406],"genre_scores_gemma":[0.3413337,0.0013024636,0.6185994,0.0015480993,0.00034040812,0.00051690295,0.027275126,0.00053411745,0.008549731],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977278,0.0008541748,0.00014199894,0.00073141477,0.00039392387,0.0001506712],"domain_scores_gemma":[0.99565154,0.0026157598,0.00017840162,0.0005450776,0.0008268937,0.00018230236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031831039,0.0014346296,0.00096494285,0.0032387176,0.0005991137,0.001942718,0.001727281,0.0022948724,0.006930813],"category_scores_gemma":[0.012183299,0.00034959055,0.0017659519,0.0015407709,0.0005902188,0.004615306,0.003040547,0.0019718346,0.004292647],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004938825,0.00053139974,0.005069228,0.0010366864,0.00024869488,0.00039136634,0.0011984515,0.027190313,0.03158103,0.009342389,0.03525458,0.887662],"study_design_scores_gemma":[0.00013608123,0.0004369794,0.0063829217,0.00021189029,0.0003136802,0.0008130235,0.00088237715,0.8172116,0.037340574,0.067767486,0.06837072,0.00013258417],"about_ca_topic_score_codex":0.007581793,"about_ca_topic_score_gemma":0.009883882,"teacher_disagreement_score":0.007581793,"about_ca_system_score_codex":0.0011105901,"about_ca_system_score_gemma":0.0021921154,"threshold_uncertainty_score":0.02318585},"labels":[],"label_agreement":null},{"id":"W4405702407","doi":"10.54660/.ijmrge.2024.5.6.1279-1286","title":"Analyzing and Mitigating Dataset Artifacts in Natural Language Inference Models Using ELECTRA","year":2024,"lang":"en","type":"article","venue":"International Journal of Multidisciplinary Research and Growth Evaluation","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Inference; Computer science; Artificial intelligence; Robustness (evolution); Weighting; Regularization (linguistics); Adversarial system; Generalization; Natural language understanding; Artifact (error); Machine learning; Natural language processing; Natural language; Mathematics","score_opus":0.14799849305238508,"score_gpt":0.4685742009250235,"score_spread":0.32057570787263845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405702407","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062977254,0.00093622104,0.92892,0.0013758462,0.00013293503,0.00021135151,0.00078101497,0.0031728195,0.0014925982],"genre_scores_gemma":[0.5935943,0.0007229223,0.39551285,0.0009777036,0.00017389114,0.00055438577,0.005176858,0.00066887256,0.0026182428],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950771,0.0028485826,0.00025908012,0.000946841,0.0007102117,0.00015827724],"domain_scores_gemma":[0.9599428,0.0278945,0.0016005194,0.008108381,0.0020746521,0.0003792035],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015151021,0.001645155,0.0010834628,0.0013674521,0.0008336135,0.002036965,0.0028572378,0.0014802578,0.0015004699],"category_scores_gemma":[0.051179487,0.0007218035,0.0012836735,0.0011557009,0.0019108587,0.0052048056,0.005451712,0.0045956345,0.0008408594],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005628761,0.00041918553,0.019866193,0.0005382184,0.00047448103,0.0004692482,0.00072815566,0.66014606,0.00957117,0.031639654,0.010067809,0.26551694],"study_design_scores_gemma":[0.000020977192,0.00009141898,0.00070137996,0.000029461098,0.000043558248,0.00012080166,0.000067900546,0.97226703,0.0037333888,0.020411728,0.0024970677,0.000015293212],"about_ca_topic_score_codex":0.0024579265,"about_ca_topic_score_gemma":0.004876954,"teacher_disagreement_score":0.984849,"about_ca_system_score_codex":0.0010274017,"about_ca_system_score_gemma":0.0017043874,"threshold_uncertainty_score":0.08012724},"labels":[],"label_agreement":null},{"id":"W4405727497","doi":"10.1111/vco.13035","title":"Precision in Parsing: Evaluation of an Open‐Source Named Entity Recognizer (<scp>NER</scp>) in Veterinary Oncology","year":2024,"lang":"en","type":"article","venue":"Veterinary and Comparative Oncology","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of Guelph; Oakville-Trafalgar Memorial Hospital; Health Sciences Centre; Sunnybrook Health Science Centre","funders":"","keywords":"Jaccard index; Named-entity recognition; F1 score; Precision and recall; Medicine; Computer science; Veterinary medicine; Artificial intelligence","score_opus":0.2758151808341951,"score_gpt":0.44483196899707067,"score_spread":0.16901678816287558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405727497","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5338311,0.011306106,0.25793535,0.0041170935,0.0024035908,0.0020792189,0.03114495,0.12527576,0.031906907],"genre_scores_gemma":[0.59720874,0.0021583678,0.32467097,0.0014469847,0.0003861718,0.00067867304,0.05942224,0.0049195345,0.009108316],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9866006,0.0055602966,0.0015128125,0.0039847475,0.0018417775,0.00049971807],"domain_scores_gemma":[0.93701226,0.04781562,0.0016343994,0.0049444833,0.0079691345,0.00062401016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02398724,0.0021048316,0.0015735653,0.0045146397,0.0017649125,0.004151228,0.0026528009,0.0037990853,0.0035146184],"category_scores_gemma":[0.059016358,0.0007598426,0.0017977374,0.0026543795,0.0011121919,0.0060209674,0.0035757823,0.001960436,0.0049228794],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0043087546,0.0014722089,0.04825588,0.0043181484,0.0019188777,0.0017579072,0.004870414,0.06441601,0.027781246,0.0046122903,0.10004123,0.7362471],"study_design_scores_gemma":[0.001019066,0.0024904895,0.09520957,0.0014691072,0.0026403272,0.0042749853,0.004638851,0.61540955,0.13592687,0.017627846,0.11844457,0.0008488257],"about_ca_topic_score_codex":0.01350959,"about_ca_topic_score_gemma":0.01422165,"teacher_disagreement_score":0.02398724,"about_ca_system_score_codex":0.0016467965,"about_ca_system_score_gemma":0.002561525,"threshold_uncertainty_score":0.12685812},"labels":[],"label_agreement":null},{"id":"W4405767395","doi":"10.48550/arxiv.2412.16971","title":"Part-Of-Speech Sensitivity of Routers in Mixture of Experts Models","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Grand Équipement National De Calcul Intensif","keywords":"Security token; Computer science; Routing (electronic design automation); Natural language processing; Artificial intelligence; Computer network","score_opus":0.08469193516844202,"score_gpt":0.19602315259198916,"score_spread":0.11133121742354714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405767395","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6702438,0.0005825958,0.32423723,0.0009112836,0.00006583144,0.000067641966,0.00025161143,0.0005536914,0.0030863096],"genre_scores_gemma":[0.98044765,0.00019062634,0.015417546,0.00013963196,0.000033359443,0.00005055129,0.0002968266,0.00009792691,0.003325859],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983974,0.00067409524,0.00004785376,0.00050824205,0.0001272202,0.00024519768],"domain_scores_gemma":[0.98527414,0.011613551,0.000890714,0.00087972666,0.00082425436,0.00051752577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004947753,0.0010140177,0.0012712985,0.001152128,0.0006277493,0.0021649315,0.0016925959,0.002125392,0.0027103801],"category_scores_gemma":[0.024867391,0.000944885,0.001034414,0.0006255786,0.0014270548,0.0043245647,0.0016103704,0.0028777104,0.0008073693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010503244,0.00013652805,0.019508785,0.000117215124,0.00024945798,0.00035615795,0.001232655,0.90667504,0.0064594885,0.03740111,0.0019674315,0.024845783],"study_design_scores_gemma":[0.000011286657,0.000042465042,0.0014347556,0.000010775321,0.000026799507,0.00006084774,0.000087229215,0.98287946,0.0008839724,0.014307248,0.00023365865,0.00002161155],"about_ca_topic_score_codex":0.008795805,"about_ca_topic_score_gemma":0.005634811,"teacher_disagreement_score":0.008795805,"about_ca_system_score_codex":0.0017196253,"about_ca_system_score_gemma":0.0007537025,"threshold_uncertainty_score":0.026166499},"labels":[],"label_agreement":null},{"id":"W4405828709","doi":"10.1145/3704440.3704777","title":"Comparative Analysis and Optimization of LoRA Adapter Co-serving for Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Adapter (computing); Computer science; Operating system","score_opus":0.04261131942652414,"score_gpt":0.313638369113908,"score_spread":0.2710270496873839,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405828709","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8111968,0.007157963,0.12592486,0.0014858486,0.0004465549,0.00035426748,0.0013597503,0.026176067,0.025897807],"genre_scores_gemma":[0.93160504,0.0007619129,0.059816822,0.0002561221,0.00006772432,0.00018656945,0.002631228,0.0013917119,0.0032827833],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960907,0.0015130658,0.00022338526,0.0005098634,0.00069056684,0.0009723227],"domain_scores_gemma":[0.9824913,0.011313603,0.00041362978,0.0027473788,0.002304191,0.0007299106],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052106306,0.0017346744,0.0013661348,0.0014477865,0.0010534835,0.0022183983,0.0027584436,0.0013180408,0.007027909],"category_scores_gemma":[0.023947323,0.0005518445,0.0010847569,0.0018132472,0.00079486636,0.003414458,0.0019229056,0.0018350335,0.002317124],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036643543,0.0016679785,0.01020668,0.0006268965,0.00038780327,0.00032172384,0.00038386983,0.64245045,0.010968094,0.00870955,0.026551625,0.29406103],"study_design_scores_gemma":[0.00009787199,0.00024317457,0.0010781841,0.00001627016,0.00007337842,0.00006180249,0.00021474902,0.9920596,0.0026895134,0.0015041614,0.001942396,0.000018799237],"about_ca_topic_score_codex":0.015356161,"about_ca_topic_score_gemma":0.019748092,"teacher_disagreement_score":0.015356161,"about_ca_system_score_codex":0.0026041723,"about_ca_system_score_gemma":0.0034451298,"threshold_uncertainty_score":0.030533552},"labels":[],"label_agreement":null},{"id":"W4405891061","doi":"10.1002/9781394287024.ch2","title":"Natural Language Processing in Healthcare","year":2024,"lang":"en","type":"other","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Université du Québec à Montréal","funders":"","keywords":"Health care; Computer science; Political science","score_opus":0.01490392371839967,"score_gpt":0.3028292891516681,"score_spread":0.28792536543326847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405891061","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070392443,0.077107,0.4382945,0.078675866,0.0042937584,0.0006999204,0.004634142,0.0039927945,0.38526264],"genre_scores_gemma":[0.24700287,0.10626976,0.45847812,0.019764325,0.006415011,0.0013529185,0.009494861,0.0012366818,0.14998549],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99723035,0.0012894605,0.00024208103,0.0004046833,0.0007184287,0.0001150277],"domain_scores_gemma":[0.99609035,0.0026267227,0.00020627011,0.0004358359,0.0005242495,0.00011655762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003182756,0.00055758003,0.000534935,0.0022745857,0.0011208663,0.006045636,0.0009590937,0.002152358,0.020875657],"category_scores_gemma":[0.008803915,0.0002967533,0.0006345378,0.0030713668,0.0022610377,0.0051978827,0.0023275686,0.0018358409,0.008998737],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004389751,0.000053300497,0.00095671095,0.0011065917,0.00004178552,0.00034217638,0.00079748454,0.004366397,0.0014321712,0.39441153,0.12871939,0.46772853],"study_design_scores_gemma":[0.000011794198,0.00002637469,0.00084203284,0.00066044315,0.000014974125,0.00049619854,0.00067934935,0.011081917,0.0010663684,0.46935117,0.51573247,0.00003694761],"about_ca_topic_score_codex":0.0028074887,"about_ca_topic_score_gemma":0.0017730851,"teacher_disagreement_score":0.020875657,"about_ca_system_score_codex":0.001848514,"about_ca_system_score_gemma":0.0026968431,"threshold_uncertainty_score":0.06983602},"labels":[],"label_agreement":null},{"id":"W4405903515","doi":"10.48550/arxiv.2412.19726","title":"Position: Theory of Mind Benchmarks are Broken for Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Canada Excellence Research Chairs, Government of Canada","keywords":"Context (archaeology); Computer science; Cognitive science; Psychology; History; Archaeology","score_opus":0.05527061287191093,"score_gpt":0.1994037170977207,"score_spread":0.14413310422580977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405903515","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13604775,0.0065228045,0.64091945,0.099285506,0.0035724367,0.00037485795,0.0017506977,0.0084229745,0.103103526],"genre_scores_gemma":[0.82804024,0.0013212113,0.15078147,0.009858605,0.0014181276,0.0007419768,0.0017330894,0.0023629994,0.003742206],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96349436,0.013121083,0.0023398723,0.007059865,0.012384826,0.0015999728],"domain_scores_gemma":[0.7615302,0.14157806,0.014010376,0.05433262,0.020921707,0.0076270066],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04652021,0.0020602068,0.0022663083,0.004289959,0.002696388,0.012001813,0.0060452814,0.0052425256,0.010399143],"category_scores_gemma":[0.24497613,0.0009767928,0.0019277004,0.0020610755,0.01239145,0.03302665,0.009668084,0.013558341,0.0030925586],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039407818,0.00020873682,0.007670311,0.0005488291,0.0002133576,0.000093075876,0.0013729708,0.007866264,0.0011985089,0.86303645,0.020065434,0.09733191],"study_design_scores_gemma":[0.00006897913,0.00026277668,0.0023003144,0.0003969187,0.00006242208,0.00017746698,0.0005323139,0.031140776,0.0029361262,0.93779874,0.024214305,0.000108860004],"about_ca_topic_score_codex":0.002719164,"about_ca_topic_score_gemma":0.0013805164,"teacher_disagreement_score":0.95347977,"about_ca_system_score_codex":0.004258548,"about_ca_system_score_gemma":0.0036039045,"threshold_uncertainty_score":0.24602538},"labels":[],"label_agreement":null},{"id":"W4405968021","doi":"10.1109/tse.2024.3519464","title":"<i>Look Before You Leap:</i> An Exploratory Study of Uncertainty Analysis for Large Language Models","year":2025,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"JST-Mirai Program; Japan Science and Technology Agency; Japan Society for the Promotion of Science","keywords":"Computer science; Exploratory analysis; Programming language; Exploratory research; Data science; Software engineering","score_opus":0.016702405548157565,"score_gpt":0.25202852066838805,"score_spread":0.2353261151202305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405968021","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16837496,0.0045240126,0.7834773,0.027589554,0.000388861,0.00050498545,0.00110629,0.0017325219,0.012301472],"genre_scores_gemma":[0.74407494,0.0017816476,0.24380971,0.0042411126,0.00041619907,0.00067152,0.0013552102,0.000761191,0.0028884236],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98900235,0.0070697316,0.00033921088,0.0010455747,0.002293453,0.00024971727],"domain_scores_gemma":[0.8517069,0.13025579,0.0034802365,0.0076790503,0.0059913113,0.000886794],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015577175,0.0010389036,0.0007150014,0.0019274419,0.0013155438,0.004187532,0.002055027,0.0021119134,0.003232282],"category_scores_gemma":[0.11879003,0.00060782215,0.0013239123,0.0015278261,0.0037649986,0.011775701,0.0034539783,0.0063892226,0.0007567147],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000904364,0.0005586096,0.059254542,0.0021739288,0.0009347508,0.0016516044,0.021533696,0.1159967,0.019598255,0.25719944,0.044190105,0.47600406],"study_design_scores_gemma":[0.00008329784,0.00068320805,0.02248672,0.0010278174,0.0001764455,0.0012455018,0.0066548786,0.5717868,0.020463673,0.30791423,0.06708784,0.00038953626],"about_ca_topic_score_codex":0.0061615766,"about_ca_topic_score_gemma":0.006248746,"teacher_disagreement_score":0.015577175,"about_ca_system_score_codex":0.0019725198,"about_ca_system_score_gemma":0.0016442703,"threshold_uncertainty_score":0.08238101},"labels":[],"label_agreement":null},{"id":"W4405995766","doi":"10.1080/09540091.2024.2445249","title":"Towards enhanced assessment question classification: a study using machine learning, deep learning, and generative AI","year":2025,"lang":"en","type":"article","venue":"Connection Science","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Quest University Canada","funders":"","keywords":"Computer science; Artificial intelligence; Generative grammar; Machine learning; Deep learning","score_opus":0.041588579686941005,"score_gpt":0.37815509574483697,"score_spread":0.33656651605789595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405995766","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.786058,0.009289071,0.17472044,0.0031974136,0.0005169555,0.0009051981,0.0021541275,0.0039903345,0.019168511],"genre_scores_gemma":[0.93526614,0.0008720702,0.056095872,0.00039083714,0.00014005983,0.00018106999,0.0029248435,0.00009556507,0.0040336247],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945391,0.0030452616,0.00028796768,0.0009850862,0.0009326791,0.00020992715],"domain_scores_gemma":[0.97801125,0.015159004,0.0011126727,0.0018790496,0.002962911,0.0008749541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007063384,0.001330484,0.00086848025,0.0035538191,0.00059773825,0.0031283975,0.0017436775,0.0015446204,0.0021662826],"category_scores_gemma":[0.026520269,0.0002978596,0.0009890627,0.0025072908,0.00059776707,0.005555568,0.0016548242,0.0031005752,0.0015876066],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010845914,0.0018542921,0.10817518,0.0012707759,0.00032921004,0.0002591971,0.002125032,0.04579453,0.006914102,0.007867143,0.011917946,0.812408],"study_design_scores_gemma":[0.00010799789,0.00077179563,0.03919765,0.00023710792,0.00018002233,0.0003122148,0.0015801294,0.9133257,0.012654352,0.012216834,0.019312523,0.000103638624],"about_ca_topic_score_codex":0.0076895826,"about_ca_topic_score_gemma":0.007135124,"teacher_disagreement_score":0.0076895826,"about_ca_system_score_codex":0.0024956923,"about_ca_system_score_gemma":0.0015621716,"threshold_uncertainty_score":0.037355244},"labels":[],"label_agreement":null},{"id":"W4406040593","doi":"10.62051/h08exg91","title":"Comparative Evaluation of GPT, BERT, and XLNet: Insights into Their Performance and Applicability in NLP Tasks","year":2024,"lang":"en","type":"article","venue":"Transactions on Computer Science and Intelligent Systems Research","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Transformer; Language model; Artificial intelligence; Natural language processing; Natural language understanding; Encoder; Machine learning; Comprehension; Generative grammar; Natural language","score_opus":0.16900054709282167,"score_gpt":0.3991781544185883,"score_spread":0.23017760732576661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406040593","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6901668,0.011028479,0.24146698,0.004016522,0.0010404666,0.00081530976,0.0043313145,0.019761503,0.027372662],"genre_scores_gemma":[0.89839625,0.0022435328,0.08385669,0.0006796234,0.00010871755,0.00052234106,0.007488951,0.00058153836,0.006122295],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998495,0.0006419175,0.00009945356,0.00038596292,0.00023817658,0.00013939658],"domain_scores_gemma":[0.99429095,0.0041155964,0.00018785964,0.0005622194,0.0005839545,0.0002593264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039872266,0.0022506842,0.00096692337,0.0011983509,0.0004355957,0.0013124545,0.002141664,0.0018737591,0.0026669875],"category_scores_gemma":[0.014133565,0.00051116705,0.00073574454,0.0009076362,0.00083422975,0.0039678654,0.0017201754,0.0031459075,0.001261056],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014123489,0.0007623179,0.011996611,0.00087706203,0.00040197314,0.0002694023,0.0004489786,0.5429724,0.003945759,0.0053434535,0.0148613565,0.41670835],"study_design_scores_gemma":[0.00006926663,0.0006010209,0.002195559,0.00009073998,0.0000969796,0.00008985876,0.00018431124,0.9860373,0.004367119,0.003747972,0.0024826007,0.000037315123],"about_ca_topic_score_codex":0.012688984,"about_ca_topic_score_gemma":0.019733567,"teacher_disagreement_score":0.012688984,"about_ca_system_score_codex":0.0017730285,"about_ca_system_score_gemma":0.0017109497,"threshold_uncertainty_score":0.025230289},"labels":[],"label_agreement":null},{"id":"W4406094712","doi":"10.1080/00330124.2024.2434455","title":"Comparing the Spatial Querying Capacity of Large Language Models: OpenAI’s ChatGPT and Google’s Gemini Pro","year":2025,"lang":"en","type":"article","venue":"The Professional Geographer","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Geography; World Wide Web","score_opus":0.03603262636662022,"score_gpt":0.29434057528629387,"score_spread":0.2583079489196737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406094712","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8994257,0.0021203812,0.048739653,0.0023790442,0.0004846755,0.0005504623,0.0068917433,0.023562314,0.01584612],"genre_scores_gemma":[0.94330186,0.00039905604,0.04223089,0.0004885079,0.00010044813,0.0002980122,0.010072191,0.00071905233,0.0023899549],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9914904,0.005040954,0.0006409065,0.001103172,0.0013670214,0.00035759422],"domain_scores_gemma":[0.941658,0.048942473,0.00080973865,0.004626404,0.0028473204,0.0011160352],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010986669,0.0013824728,0.0011872081,0.0021938696,0.0008722574,0.0029061134,0.002864101,0.001958874,0.0026134541],"category_scores_gemma":[0.05454201,0.000633936,0.001058016,0.002149423,0.0013633546,0.008461502,0.004308305,0.002081732,0.0013250433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.014207084,0.003141837,0.093166724,0.00460625,0.0014168555,0.0018100931,0.025118556,0.30829048,0.015817309,0.03479642,0.09854064,0.3990878],"study_design_scores_gemma":[0.00040537879,0.0007386,0.009729183,0.00011502842,0.00020387069,0.00035269564,0.0037855557,0.9532375,0.0058159167,0.011095436,0.014312659,0.00020814441],"about_ca_topic_score_codex":0.043655463,"about_ca_topic_score_gemma":0.03585339,"teacher_disagreement_score":0.043655463,"about_ca_system_score_codex":0.0020419892,"about_ca_system_score_gemma":0.0018581456,"threshold_uncertainty_score":0.08680272},"labels":[],"label_agreement":null},{"id":"W4406248884","doi":"10.1075/ml.24027.deg","title":"NLP and education","year":2024,"lang":"en","type":"article","venue":"The Mental Lexicon","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"","keywords":"Natural language processing; Semantic similarity; Computer science; Artificial intelligence; Similarity (geometry); Reading comprehension; Test (biology); Comprehension; Scale (ratio); Cloze test; Portuguese; Brazilian Portuguese; Reading (process); Information retrieval; Linguistics","score_opus":0.016931853376873304,"score_gpt":0.27242788788490674,"score_spread":0.25549603450803343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406248884","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07523586,0.012784446,0.1191552,0.04114472,0.0010369449,0.00021285708,0.0017774227,0.0012884149,0.7473641],"genre_scores_gemma":[0.9088818,0.0057899114,0.02762946,0.0020369955,0.00029897867,0.00015333509,0.00090985507,0.000258485,0.05404125],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946234,0.0030430753,0.00029821813,0.00067339215,0.0011637871,0.00019816583],"domain_scores_gemma":[0.9851555,0.009680389,0.0012169518,0.0019725133,0.0015338431,0.00044085478],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037675588,0.00039706414,0.0003295584,0.0024900225,0.0011801095,0.0076950463,0.0005861938,0.0010132659,0.025876427],"category_scores_gemma":[0.0254062,0.00015020749,0.00024928266,0.003704839,0.0028953142,0.0062681646,0.003416394,0.0013652386,0.004314776],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037283407,0.000083034734,0.012403652,0.00045334306,0.00001846497,0.00046998594,0.008421002,0.0014420939,0.00073713757,0.3939525,0.022571791,0.5594097],"study_design_scores_gemma":[0.000014590783,0.00006298745,0.015692852,0.0010172408,0.000017130493,0.0011795352,0.009830721,0.0058318516,0.001386771,0.32501072,0.6399255,0.000030061052],"about_ca_topic_score_codex":0.0032555086,"about_ca_topic_score_gemma":0.002395982,"teacher_disagreement_score":0.025876427,"about_ca_system_score_codex":0.002292749,"about_ca_system_score_gemma":0.0028743795,"threshold_uncertainty_score":0.086565316},"labels":[],"label_agreement":null},{"id":"W4406457848","doi":"10.1109/bigdata62323.2024.10825619","title":"From Graph Paths to Natural Language: Enhancing LLM Reasoning for Multi-choice Question-Answering Tasks","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Question answering; Computer science; Natural language; Graph; Artificial intelligence; Natural language processing; Theoretical computer science","score_opus":0.01681412233054059,"score_gpt":0.3105731279335161,"score_spread":0.2937590056029755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406457848","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020871853,0.0004170099,0.9663293,0.0010948512,0.000046850797,0.00024875117,0.0010378639,0.008177176,0.0017763252],"genre_scores_gemma":[0.2457099,0.00033123285,0.74825126,0.00039124055,0.000049915812,0.0002681707,0.0031262357,0.0003279385,0.0015440616],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99831414,0.0007853823,0.000097117125,0.00046384995,0.0002464934,0.00009304635],"domain_scores_gemma":[0.9968264,0.002179157,0.00016743252,0.00044067475,0.0002730524,0.00011323606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001999849,0.0010114064,0.00051920116,0.0021633066,0.00055010227,0.0012836336,0.0016536926,0.0012521913,0.004478403],"category_scores_gemma":[0.009327657,0.00045676308,0.0017102364,0.0015616876,0.00086233986,0.0059719426,0.002846968,0.0019021002,0.0011320524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004979974,0.000642437,0.005585887,0.0010376434,0.00021481811,0.00067014596,0.0025321331,0.09348746,0.020564897,0.07360602,0.02293148,0.77822906],"study_design_scores_gemma":[0.00007357198,0.00009188591,0.0008022304,0.000064898384,0.0000935349,0.00016134263,0.00030793427,0.8401946,0.009522557,0.13257265,0.016068814,0.000046006113],"about_ca_topic_score_codex":0.01064615,"about_ca_topic_score_gemma":0.017056562,"teacher_disagreement_score":0.01064615,"about_ca_system_score_codex":0.0011751328,"about_ca_system_score_gemma":0.0017688071,"threshold_uncertainty_score":0.02116841},"labels":[],"label_agreement":null},{"id":"W4406457934","doi":"10.1109/bigdata62323.2024.10825103","title":"SoK: Prompt Hacking of Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Hacker; Computer science; Computer security","score_opus":0.026333818958019957,"score_gpt":0.28419054310251796,"score_spread":0.257856724144498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406457934","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17252651,0.0010186856,0.76783967,0.0034344038,0.0003991603,0.0006033204,0.0022263299,0.04579364,0.0061582425],"genre_scores_gemma":[0.80333817,0.00041505773,0.18842953,0.0009786688,0.00019832412,0.00028904245,0.0024705057,0.0014995948,0.0023811064],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9941843,0.0028947035,0.0003472129,0.00084012264,0.0013986783,0.0003348518],"domain_scores_gemma":[0.9629469,0.025740733,0.0028402824,0.0056943665,0.0021314821,0.0006462724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005479563,0.0011272862,0.0007228643,0.0017349627,0.0007264702,0.0026853203,0.0014989304,0.0016548652,0.002147974],"category_scores_gemma":[0.044168174,0.00041462036,0.0010814654,0.00068435527,0.0014274686,0.005653043,0.0033208337,0.0028006288,0.0009259797],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021949103,0.00093181257,0.07920634,0.0022196793,0.0008107143,0.0022851324,0.009477583,0.17022109,0.05733136,0.10017573,0.0493076,0.52583796],"study_design_scores_gemma":[0.000057325866,0.00021488254,0.004105126,0.00010057291,0.00009304978,0.0005432053,0.00072424195,0.90179044,0.018093826,0.063776985,0.010394406,0.000105908235],"about_ca_topic_score_codex":0.00222011,"about_ca_topic_score_gemma":0.0023761673,"teacher_disagreement_score":0.005479563,"about_ca_system_score_codex":0.0009173877,"about_ca_system_score_gemma":0.0015863661,"threshold_uncertainty_score":0.028979003},"labels":[],"label_agreement":null},{"id":"W4406711870","doi":"10.1007/s10115-025-02338-0","title":"Unlocking wisdom: enhancing biomedical question answering with domain knowledge","year":2025,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Question answering; Domain (mathematical analysis); Computer science; Information retrieval; Domain knowledge; Data science; Knowledge management; Mathematics","score_opus":0.007620119195518866,"score_gpt":0.24825661733302845,"score_spread":0.2406364981375096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406711870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05430708,0.0055256314,0.91614616,0.008659903,0.00047409607,0.00035806975,0.0023113918,0.0053603905,0.0068572927],"genre_scores_gemma":[0.39431584,0.0025881,0.59233534,0.001904283,0.0007463148,0.00024993243,0.0047643883,0.0004832764,0.002612555],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928377,0.0033607176,0.00058957713,0.0015119177,0.0014720757,0.00022796869],"domain_scores_gemma":[0.95061433,0.03944937,0.0013072796,0.004724205,0.0031408677,0.0007638964],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009022497,0.0012690555,0.0016889955,0.0044372943,0.0011763973,0.004825445,0.0025536388,0.0027782829,0.0036089814],"category_scores_gemma":[0.059460215,0.00075652974,0.0017655015,0.0029734473,0.0012321578,0.01377882,0.005940647,0.0038271116,0.002043953],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006706817,0.0009689598,0.010387246,0.0016118119,0.00047161267,0.00033980847,0.0030128611,0.028899025,0.01082335,0.025925258,0.023621384,0.89326817],"study_design_scores_gemma":[0.00013924374,0.00025613408,0.0029702797,0.0003355075,0.0006850734,0.00042304225,0.0014282042,0.6020393,0.011197248,0.34973451,0.03066989,0.00012160852],"about_ca_topic_score_codex":0.0025941692,"about_ca_topic_score_gemma":0.0048789214,"teacher_disagreement_score":0.009022497,"about_ca_system_score_codex":0.00090443034,"about_ca_system_score_gemma":0.0020635496,"threshold_uncertainty_score":0.04771608},"labels":[],"label_agreement":null},{"id":"W4406728115","doi":"10.1109/mnet.2025.3532857","title":"PrismPrompt: Layering Prompt-Enhanced Cloud-Edge Collaborative Language Model Toward Healthcare","year":2025,"lang":"en","type":"article","venue":"IEEE Network","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Layering; Cloud computing; Computer science; Enhanced Data Rates for GSM Evolution; Health care; Telecommunications; Operating system","score_opus":0.02184038712467231,"score_gpt":0.2899834405131377,"score_spread":0.2681430533884654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406728115","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005749592,0.00024279398,0.9774699,0.00077346974,0.00008814296,0.0002062105,0.0006878471,0.0130549725,0.0017271376],"genre_scores_gemma":[0.20762268,0.0006014241,0.7819858,0.0011601618,0.00011672384,0.00041949627,0.002891149,0.0012152514,0.003987294],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978416,0.0008642481,0.0002139868,0.00044934134,0.0005096317,0.000121184676],"domain_scores_gemma":[0.9967709,0.0015675597,0.00019934606,0.00076009263,0.00044279906,0.0002591289],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027333458,0.0010121043,0.0007737764,0.0007646085,0.0006250148,0.0027853637,0.002198032,0.0012880266,0.004181299],"category_scores_gemma":[0.011823545,0.00048641488,0.0015175843,0.0007736823,0.00064883934,0.0041622263,0.0044869,0.002273281,0.0025340118],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016131279,0.00078532036,0.008284861,0.0012051441,0.00035381303,0.001577604,0.0023698842,0.25821257,0.024531856,0.092881136,0.064753726,0.5434309],"study_design_scores_gemma":[0.00008365272,0.000105757645,0.00025819318,0.000051148858,0.000057419646,0.00025366552,0.00017163753,0.914895,0.008200955,0.04879566,0.027073098,0.00005392266],"about_ca_topic_score_codex":0.0053982497,"about_ca_topic_score_gemma":0.008687957,"teacher_disagreement_score":0.0053982497,"about_ca_system_score_codex":0.0009947688,"about_ca_system_score_gemma":0.0042752316,"threshold_uncertainty_score":0.014455497},"labels":[],"label_agreement":null},{"id":"W4406753387","doi":"10.1007/978-3-031-78548-1_25","title":"PRAGyan - Connecting the Dots in Tweets","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; World Wide Web; Artificial intelligence; Information retrieval; Natural language processing","score_opus":0.021017006780768368,"score_gpt":0.2578272913690574,"score_spread":0.23681028458828904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406753387","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04899133,0.0014683274,0.7825522,0.010846957,0.0029655935,0.0005026995,0.004742803,0.0120423725,0.13588767],"genre_scores_gemma":[0.37762946,0.0013359432,0.48298594,0.0020591759,0.0005731919,0.00062902557,0.0058189365,0.0034347952,0.1255335],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988826,0.00043808378,0.00004282902,0.00032205964,0.00021162741,0.000102808124],"domain_scores_gemma":[0.99679023,0.0016084752,0.00014380728,0.00086906797,0.0004288991,0.00015957364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015570289,0.0006967072,0.00066407473,0.0011939041,0.0025460483,0.0033314514,0.00106071,0.0012198066,0.031364825],"category_scores_gemma":[0.00972217,0.000677578,0.00090599194,0.002243488,0.0010599016,0.008757541,0.005186356,0.002714655,0.015324583],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054477283,0.00014856986,0.0040524467,0.00056743925,0.00012454043,0.00034936707,0.0037131475,0.011987098,0.008331484,0.46378222,0.10786403,0.3985349],"study_design_scores_gemma":[0.00006372649,0.0001597618,0.0020332558,0.00025446125,0.00012105359,0.00033768394,0.0023906531,0.13931152,0.012238212,0.489259,0.35374647,0.000084177555],"about_ca_topic_score_codex":0.002402981,"about_ca_topic_score_gemma":0.0038474218,"teacher_disagreement_score":0.031364825,"about_ca_system_score_codex":0.0007858498,"about_ca_system_score_gemma":0.0011119748,"threshold_uncertainty_score":0.10492581},"labels":[],"label_agreement":null},{"id":"W4406771321","doi":"10.1007/978-981-96-1024-2_8","title":"Evaluation of Retrieval-Augmented Generation: A Survey","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":104,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Information retrieval; Computer science; Thesaurus; Natural language processing","score_opus":0.17110464924601465,"score_gpt":0.3563433662537061,"score_spread":0.18523871700769146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406771321","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06901086,0.5157677,0.3466562,0.0025217975,0.0013184097,0.0011651852,0.005405778,0.016317662,0.041836366],"genre_scores_gemma":[0.35876328,0.1611558,0.43092388,0.0021611066,0.0015041908,0.00058156543,0.020971132,0.0029332915,0.021005778],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98868424,0.0048718043,0.0007662502,0.0012494581,0.0041009313,0.00032736862],"domain_scores_gemma":[0.9824306,0.011915381,0.00045995004,0.002002999,0.0029544167,0.00023674552],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009622417,0.0019830775,0.002797654,0.0045693615,0.00059196865,0.0029203186,0.0041139266,0.0018760631,0.011374844],"category_scores_gemma":[0.0281756,0.0006052185,0.0012788854,0.004198235,0.0006855403,0.0035699797,0.0015731809,0.000852401,0.004131423],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067363796,0.0003159175,0.0016560346,0.0016330671,0.00018260135,0.00002691285,0.00005877993,0.007747243,0.0017319991,0.0012903424,0.012984081,0.9716994],"study_design_scores_gemma":[0.0012156109,0.007127576,0.023370767,0.00378355,0.0028540103,0.0030087233,0.0013151224,0.6197556,0.06587233,0.022074696,0.24916053,0.0004615056],"about_ca_topic_score_codex":0.008389433,"about_ca_topic_score_gemma":0.007567773,"teacher_disagreement_score":0.011374844,"about_ca_system_score_codex":0.0012096139,"about_ca_system_score_gemma":0.0021953585,"threshold_uncertainty_score":0.050888777},"labels":[],"label_agreement":null},{"id":"W4406775315","doi":"10.1007/s00521-024-10953-1","title":"AFuNet: an attention-based fusion network to classify texts in a resource-constrained language","year":2025,"lang":"en","type":"article","venue":"Neural Computing and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Athabasca University","funders":"","keywords":"Computational Science and Engineering; Computer science; Artificial intelligence; Fusion; Resource (disambiguation); Natural language processing; Artificial neural network; Machine learning; Linguistics; Philosophy","score_opus":0.013447201604617062,"score_gpt":0.2839591969143264,"score_spread":0.2705119953097094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406775315","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10518204,0.0058754473,0.8555806,0.0013892588,0.00095800054,0.00042294103,0.003920233,0.020194484,0.0064769587],"genre_scores_gemma":[0.5927027,0.00193604,0.37046486,0.0010688608,0.0009113401,0.00049930863,0.010793762,0.00085418305,0.02076897],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990181,0.00022429615,0.00006052546,0.0003210566,0.0002316552,0.00014445715],"domain_scores_gemma":[0.9986286,0.0006523921,0.000076103184,0.00013364553,0.0004095403,0.00009972541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021627746,0.0018947509,0.0013924013,0.003371286,0.0009294992,0.0012609095,0.0016627343,0.0021230835,0.0042941533],"category_scores_gemma":[0.0037983945,0.00044056994,0.0011685103,0.00228386,0.00053154916,0.003718948,0.0025166017,0.0019085048,0.0025064813],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084927137,0.00039532225,0.0023389077,0.00026866648,0.0002950268,0.000210814,0.0003742893,0.01890092,0.025484681,0.0025132596,0.023276074,0.92509276],"study_design_scores_gemma":[0.00006643571,0.00032248703,0.0020684814,0.000057859386,0.0001939647,0.00018402447,0.00016771365,0.9627497,0.015048258,0.009622718,0.009457528,0.000060823593],"about_ca_topic_score_codex":0.00813069,"about_ca_topic_score_gemma":0.012725828,"teacher_disagreement_score":0.00813069,"about_ca_system_score_codex":0.0009550619,"about_ca_system_score_gemma":0.0011273264,"threshold_uncertainty_score":0.016166747},"labels":[],"label_agreement":null},{"id":"W4406809948","doi":"10.18280/isi.300114","title":"A Unified Approach to Text Summarization: Classical, Machine Learning, and Deep Learning Methods","year":2025,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; Artificial intelligence; Deep learning; Natural language processing; Machine learning","score_opus":0.01734848913071882,"score_gpt":0.2739966536290185,"score_spread":0.25664816449829964,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406809948","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017824284,0.008184444,0.98608464,0.0011257014,0.00025929345,0.000053237894,0.00021504938,0.0007523232,0.0015430253],"genre_scores_gemma":[0.12188187,0.018824456,0.8440891,0.0010062004,0.0022347774,0.0003706416,0.0013792705,0.00043875008,0.00977494],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988558,0.00043335502,0.00012164415,0.0002470735,0.00029148595,0.000050590486],"domain_scores_gemma":[0.99858546,0.00054163317,0.00015256244,0.00022383498,0.0004424395,0.000054066342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017849758,0.0012793324,0.001138237,0.0025398135,0.0005665575,0.0023538622,0.0015456636,0.0012082764,0.0022118967],"category_scores_gemma":[0.0042092376,0.000443564,0.0011018278,0.0033449682,0.0009301091,0.0052699232,0.0013814573,0.0026013388,0.0013940266],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000097296244,0.00008632943,0.00061506435,0.0013872075,0.000189866,0.000106634965,0.00043237413,0.051926892,0.0072423033,0.07940688,0.014064371,0.84444475],"study_design_scores_gemma":[0.000037739195,0.0002677459,0.0012158899,0.00044606894,0.00015000475,0.00021617663,0.00029407916,0.66478384,0.011798432,0.24309029,0.07759532,0.00010443286],"about_ca_topic_score_codex":0.0018394224,"about_ca_topic_score_gemma":0.0024409494,"teacher_disagreement_score":0.0025398135,"about_ca_system_score_codex":0.0012652496,"about_ca_system_score_gemma":0.0010663855,"threshold_uncertainty_score":0.009440005},"labels":[],"label_agreement":null},{"id":"W4406825368","doi":"10.18280/mmep.120123","title":"Sarcasm Detection an Explainable AI Approach for Reddit Political Text","year":2025,"lang":"en","type":"article","venue":"Mathematical Modelling and Engineering Problems","topic":"Topic Modeling","field":"Computer Science","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Sarcasm; Politics; Artificial intelligence; Natural language processing; Computer science; History; Psychology; Political science; Irony; Linguistics; Philosophy","score_opus":0.025100655947616686,"score_gpt":0.2339263670081814,"score_spread":0.20882571106056472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406825368","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20559792,0.0011873636,0.75068295,0.002857039,0.00049264723,0.0010663215,0.0033792376,0.018542327,0.016194172],"genre_scores_gemma":[0.70934093,0.0005915973,0.26992282,0.00042724944,0.00021139633,0.00047204876,0.0046498263,0.00023680349,0.014147236],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994093,0.000134507,0.000051316496,0.00018585612,0.00017470527,0.000044340602],"domain_scores_gemma":[0.9986896,0.00057364133,0.00018827207,0.00010700138,0.00040448213,0.00003712774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00081499,0.0007896699,0.0004021228,0.0015318138,0.0004293173,0.001046793,0.0007767962,0.00064115477,0.0035909913],"category_scores_gemma":[0.002758656,0.00018756354,0.0007028834,0.0005651604,0.00022781118,0.0011604936,0.00048813716,0.00096230925,0.0019450133],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037769676,0.00044664185,0.015499973,0.000638077,0.00022499682,0.00057844527,0.0017681563,0.008724503,0.058143802,0.004199447,0.01986163,0.8895366],"study_design_scores_gemma":[0.000050035134,0.0005803134,0.034986988,0.00014666324,0.00022716969,0.00074014085,0.0013582638,0.86437064,0.05750404,0.011756077,0.028190205,0.00008952648],"about_ca_topic_score_codex":0.0014192585,"about_ca_topic_score_gemma":0.0022536726,"teacher_disagreement_score":0.0035909913,"about_ca_system_score_codex":0.000473993,"about_ca_system_score_gemma":0.00043757085,"threshold_uncertainty_score":0.012013078},"labels":[],"label_agreement":null},{"id":"W4406865074","doi":"10.2196/65984","title":"Large Language Model Applications for Health Information Extraction in Oncology: Scoping Review","year":2025,"lang":"en","type":"review","venue":"JMIR Cancer","topic":"Topic Modeling","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"BC Cancer Agency; University of Waterloo; University of Toronto","funders":"","keywords":"Preprint; Computer science; Data science; Medicine; World Wide Web","score_opus":0.074606718448783,"score_gpt":0.5065021944205048,"score_spread":0.43189547597172184,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406865074","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00050239306,0.99170727,0.0037352964,0.0013637969,0.00020321901,0.0007425489,0.0006339188,0.00007265898,0.0010388535],"genre_scores_gemma":[0.0067670634,0.97969353,0.009393018,0.0008123659,0.00017919563,0.0020951733,0.0008420044,0.000043491542,0.00017423282],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.98117673,0.008528094,0.0062772306,0.0011311884,0.0026473291,0.00023950715],"domain_scores_gemma":[0.73557085,0.2345575,0.012161772,0.004091359,0.013070055,0.0005484893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03979845,0.002027066,0.0045061186,0.02335501,0.0012806643,0.005248263,0.0038212256,0.003326264,0.0067250403],"category_scores_gemma":[0.20482376,0.0015155937,0.009605555,0.022250483,0.0022046168,0.006641968,0.0040315357,0.0029748331,0.0015556478],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014289918,0.00006304,0.0009969189,0.61334586,0.0022899874,0.0001588051,0.0009067745,0.0012746805,0.0002667901,0.0033936417,0.0075193965,0.36964113],"study_design_scores_gemma":[0.000054255237,0.00014109546,0.0016168308,0.90090793,0.007238428,0.00028277538,0.0006628515,0.000872555,0.00052972115,0.003871349,0.083751686,0.000070437774],"about_ca_topic_score_codex":0.010814383,"about_ca_topic_score_gemma":0.01734942,"teacher_disagreement_score":0.03979845,"about_ca_system_score_codex":0.0059867455,"about_ca_system_score_gemma":0.024640838,"threshold_uncertainty_score":0.21047688},"labels":[],"label_agreement":null},{"id":"W4406870107","doi":"10.1007/978-3-031-78952-6_44","title":"From Liberating to Questioning Tabular Data in Documents Using Knowledge Graphs","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Knowledge graph; Information retrieval; World Wide Web; Data science","score_opus":0.03582358553932706,"score_gpt":0.3071033740246556,"score_spread":0.2712797884853285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406870107","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046157283,0.0012991277,0.9634548,0.0040702163,0.0002697914,0.00008926028,0.0006307409,0.006080867,0.019489417],"genre_scores_gemma":[0.10649274,0.0038521986,0.8353827,0.0018571186,0.00030689244,0.00016942013,0.0034077934,0.004735805,0.043795377],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979943,0.00069236016,0.00015238156,0.00045526322,0.00059802405,0.00010763837],"domain_scores_gemma":[0.99098223,0.0053887162,0.0002128452,0.0025576039,0.0006667364,0.00019179696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002350294,0.00079300185,0.00084336486,0.0016453579,0.0010684915,0.00558669,0.0016596335,0.0013246496,0.01646033],"category_scores_gemma":[0.01488905,0.0008597174,0.001534376,0.0023797068,0.0027944234,0.015910752,0.0050159567,0.004308606,0.008395882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000117357195,0.000058547266,0.0005915599,0.00068940653,0.000048663547,0.0001450642,0.0028974486,0.0057064868,0.0049793352,0.3555425,0.053770795,0.5754528],"study_design_scores_gemma":[0.000020128142,0.000033481785,0.0003006222,0.0004025364,0.00006782225,0.00028568806,0.0013017799,0.037960954,0.011133173,0.6465859,0.30185792,0.000050071885],"about_ca_topic_score_codex":0.0029444993,"about_ca_topic_score_gemma":0.0032256395,"teacher_disagreement_score":0.01646033,"about_ca_system_score_codex":0.0012072283,"about_ca_system_score_gemma":0.0013141425,"threshold_uncertainty_score":0.055065274},"labels":[],"label_agreement":null},{"id":"W4406872811","doi":"10.1016/j.infsof.2025.107674","title":"XL-HQL: A HQL query generation method via XLNet and column attention","year":2025,"lang":"en","type":"article","venue":"Information and Software Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Natural Science Foundation of China","keywords":"Column (typography); Computer science; Telecommunications","score_opus":0.008798551029385699,"score_gpt":0.2539176403397022,"score_spread":0.2451190893103165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406872811","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008365748,0.0003764418,0.82348454,0.00074067595,0.00023072559,0.0005997448,0.013333313,0.1447249,0.00814392],"genre_scores_gemma":[0.14624701,0.00030064723,0.785828,0.0008541065,0.00026254109,0.0009465844,0.034989644,0.0134695,0.017101862],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984591,0.000384631,0.00014539577,0.00028184764,0.0006024146,0.0001266127],"domain_scores_gemma":[0.9971812,0.0012990426,0.00010117181,0.0004670614,0.00082479697,0.00012683254],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015477254,0.001363121,0.0008359286,0.0038777092,0.0007175516,0.0018543379,0.0020336013,0.0010451947,0.049037],"category_scores_gemma":[0.00710984,0.0007589589,0.0010515858,0.0023365438,0.0004645592,0.0030597448,0.002712789,0.0010748954,0.01371254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079005054,0.0002500697,0.004376549,0.0012154466,0.00017362423,0.00029455998,0.0006799914,0.011215152,0.024307286,0.016036171,0.26573768,0.6749234],"study_design_scores_gemma":[0.0004262543,0.00026663413,0.002278494,0.00014280631,0.00019630046,0.0003137915,0.00073522417,0.7237239,0.054650705,0.026009135,0.19111542,0.00014135918],"about_ca_topic_score_codex":0.008808743,"about_ca_topic_score_gemma":0.011934635,"teacher_disagreement_score":0.049037,"about_ca_system_score_codex":0.001021951,"about_ca_system_score_gemma":0.0014628565,"threshold_uncertainty_score":0.1640451},"labels":[],"label_agreement":null},{"id":"W4406930353","doi":"10.21203/rs.3.rs-5656576/v1","title":"RETRACTED: Repurposing the Scientific Literature with Vision-Language Models","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":true,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Columbia College","funders":"","keywords":"Repurposing; Computer science; Natural language processing; Artificial intelligence; Data science; Engineering","score_opus":0.05142588329128231,"score_gpt":0.38848649201439556,"score_spread":0.3370606087231133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406930353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04702906,0.029406678,0.6757782,0.10387334,0.030842438,0.0011193163,0.0248127,0.024392718,0.06274561],"genre_scores_gemma":[0.37058535,0.014867605,0.4455205,0.012163583,0.016999627,0.0013060732,0.053958777,0.012586177,0.07201235],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98856854,0.004584359,0.001001558,0.0021164813,0.0031574313,0.00057155185],"domain_scores_gemma":[0.9197946,0.03707022,0.0043506417,0.01617549,0.019366968,0.003242112],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.017242316,0.0019716723,0.0014701321,0.018919019,0.0035935515,0.012520768,0.005327236,0.0038233267,0.022045882],"category_scores_gemma":[0.14647086,0.0013151561,0.0033410187,0.012252506,0.0027412665,0.016773548,0.008878021,0.0067165447,0.01614054],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005718213,0.00026777847,0.0074838474,0.0021727001,0.00038891152,0.0014577807,0.005035071,0.0071719247,0.0051369746,0.054908067,0.34861383,0.5667912],"study_design_scores_gemma":[0.00014451376,0.00020374071,0.0043097558,0.001702568,0.0005504883,0.0010054627,0.005218021,0.11881825,0.008785657,0.25240204,0.60664326,0.00021635852],"about_ca_topic_score_codex":0.013943446,"about_ca_topic_score_gemma":0.013492075,"teacher_disagreement_score":0.9827577,"about_ca_system_score_codex":0.0027183986,"about_ca_system_score_gemma":0.0092603145,"threshold_uncertainty_score":0.09118718},"labels":[],"label_agreement":null},{"id":"W4406940792","doi":"10.1007/978-3-031-79029-4_29","title":"ERASMO: Leveraging Large Language Models for Enhanced Clustering Segmentation","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Cluster analysis; Segmentation; Artificial intelligence; Natural language processing","score_opus":0.020869791789530218,"score_gpt":0.27262742402814083,"score_spread":0.2517576322386106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406940792","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010625023,0.0009050617,0.925997,0.00024059496,0.00029219248,0.00015475794,0.0023802873,0.05673909,0.002666],"genre_scores_gemma":[0.09535728,0.0006004477,0.86712134,0.00055462925,0.0002696934,0.00035291785,0.013809349,0.009489469,0.012444949],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99854445,0.00033256348,0.00007664082,0.0005462338,0.00033256094,0.00016751919],"domain_scores_gemma":[0.9982622,0.0007589298,0.00007612941,0.00043091638,0.00036471648,0.00010702743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015149646,0.0023390495,0.0018871856,0.0030714339,0.001280557,0.0030859224,0.0032425865,0.0026418287,0.010257525],"category_scores_gemma":[0.0037091537,0.0013620425,0.0027784796,0.0031036483,0.0005895453,0.0034390597,0.0027898361,0.002857481,0.01669591],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008859092,0.0003796588,0.0012678246,0.00036970997,0.00045018637,0.00034833833,0.00039622773,0.06449161,0.05448545,0.006560282,0.060951315,0.80941343],"study_design_scores_gemma":[0.00004923721,0.00008137764,0.0005101349,0.000026045502,0.00009481168,0.00021135029,0.00011593593,0.9520457,0.020842545,0.012403085,0.013556961,0.00006269186],"about_ca_topic_score_codex":0.011260774,"about_ca_topic_score_gemma":0.026773866,"teacher_disagreement_score":0.011260774,"about_ca_system_score_codex":0.00092351314,"about_ca_system_score_gemma":0.0015601781,"threshold_uncertainty_score":0.03431481},"labels":[],"label_agreement":null},{"id":"W4407168991","doi":"10.1109/tlt.2025.3539104","title":"Navigating the Textual Maze: Enhancing Textual Analytical Skills Through an Innovative GAI Prompt Framework","year":2025,"lang":"en","type":"article","venue":"IEEE Transactions on Learning Technologies","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Institute for Christian Studies; University of Toronto","funders":"","keywords":"Computer science; Multimedia; Knowledge management; Natural language processing; Human–computer interaction; Artificial intelligence; Mathematics education; World Wide Web; Psychology","score_opus":0.018108583857922345,"score_gpt":0.3172257180081921,"score_spread":0.29911713415026975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407168991","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43105942,0.0006349253,0.5323659,0.0010637046,0.00018148913,0.0012906804,0.0009450018,0.01823674,0.014222081],"genre_scores_gemma":[0.62567025,0.00032998621,0.3666494,0.00029051598,0.000046620044,0.0008228354,0.00071549264,0.00039165144,0.005083309],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991702,0.00036716438,0.000044890112,0.00016948348,0.00018102591,0.00006720693],"domain_scores_gemma":[0.9926085,0.005406498,0.0004447196,0.00057503313,0.0005254675,0.00043975795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016840653,0.0008096934,0.0005275699,0.00083100545,0.0003506351,0.0014480051,0.0012911649,0.0007322552,0.005021398],"category_scores_gemma":[0.012882664,0.00021985044,0.00031094768,0.0005745597,0.0005124555,0.001782378,0.0018758872,0.0008553667,0.0014315213],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016045916,0.002083335,0.015326964,0.0021408873,0.000048277023,0.0013521904,0.014839021,0.009164918,0.10076054,0.0099785505,0.011883755,0.8308171],"study_design_scores_gemma":[0.0017492798,0.012446318,0.0897869,0.0023407035,0.00077197346,0.007362271,0.02576266,0.30767033,0.19951475,0.08020104,0.27164152,0.0007522802],"about_ca_topic_score_codex":0.00052521325,"about_ca_topic_score_gemma":0.0010257214,"teacher_disagreement_score":0.005021398,"about_ca_system_score_codex":0.00038410584,"about_ca_system_score_gemma":0.0012097463,"threshold_uncertainty_score":0.016798258},"labels":[],"label_agreement":null},{"id":"W4407187061","doi":"10.3389/frai.2025.1513674","title":"Targeted generative data augmentation for automatic metastases detection from free-text radiology reports","year":2025,"lang":"en","type":"article","venue":"Frontiers in Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"National Cancer Institute; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Context (archaeology); Task (project management); Natural language processing; Identification (biology); Artificial intelligence; F1 score; Language model; Data extraction; Machine learning; Information extraction; Information retrieval; MEDLINE","score_opus":0.0575910859937608,"score_gpt":0.3186015932620943,"score_spread":0.2610105072683335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407187061","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23684049,0.001578418,0.7238599,0.0013127216,0.0003820418,0.00028783243,0.0035707916,0.029333277,0.0028345974],"genre_scores_gemma":[0.7926858,0.00037923205,0.19056702,0.000555755,0.00012636048,0.00039515208,0.011112377,0.00057105976,0.0036072994],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991375,0.0003828029,0.00005240997,0.00023421615,0.00012852924,0.000064549335],"domain_scores_gemma":[0.99640816,0.0026364368,0.00017419527,0.00043738633,0.0002589931,0.00008475184],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013981495,0.0011709771,0.0005815019,0.00077126495,0.00023810293,0.0006136464,0.0015233526,0.00095121923,0.0022475144],"category_scores_gemma":[0.005421485,0.0004607451,0.0012670042,0.0004906499,0.0006020944,0.0009724327,0.0013229526,0.0017736516,0.0020968504],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014395135,0.0006475955,0.009811871,0.0006509676,0.00019939234,0.0010478873,0.00077501114,0.2889082,0.06192926,0.0031673135,0.018755652,0.6126674],"study_design_scores_gemma":[0.00003946895,0.0001663614,0.0015739347,0.00003109239,0.000038760518,0.0002401556,0.000088576424,0.96564364,0.025446085,0.0025350396,0.004161112,0.000035815632],"about_ca_topic_score_codex":0.003454108,"about_ca_topic_score_gemma":0.005137037,"teacher_disagreement_score":0.003454108,"about_ca_system_score_codex":0.0004846864,"about_ca_system_score_gemma":0.0009526378,"threshold_uncertainty_score":0.007518649},"labels":[],"label_agreement":null},{"id":"W4407230512","doi":"10.1016/j.eswa.2025.126648","title":"Dynamic link prediction: Using language models and graph structures for temporal knowledge graph completion with emerging entities and relations","year":2025,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Bank of Canada","funders":"","keywords":"Computer science; Knowledge graph; Graph; Link (geometry); Theoretical computer science; Artificial intelligence","score_opus":0.01790363908459883,"score_gpt":0.27997578535207374,"score_spread":0.2620721462674749,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407230512","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025645519,0.0003960007,0.9669662,0.00050860207,0.00007566496,0.00016782363,0.0015719556,0.0036938,0.00097435265],"genre_scores_gemma":[0.3905848,0.00052691344,0.5948303,0.00026224295,0.0001476579,0.00041060153,0.008790635,0.0005727187,0.0038742402],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99858594,0.0003824748,0.000097341086,0.00054216414,0.0002682583,0.0001237661],"domain_scores_gemma":[0.9934801,0.0044138706,0.00044523363,0.0007784331,0.00065808353,0.00022433362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022506574,0.0011554819,0.0011400828,0.003707346,0.0008911009,0.0017421932,0.002969951,0.001479859,0.003087375],"category_scores_gemma":[0.010813633,0.00074257725,0.0017061213,0.003151778,0.0007248493,0.004860673,0.0021067294,0.0027192314,0.0014454982],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078728254,0.0009757179,0.0060277823,0.00048304023,0.0002973238,0.0004859558,0.0007100082,0.32786298,0.0056342166,0.03187243,0.020768195,0.60409504],"study_design_scores_gemma":[0.000014267699,0.000017462124,0.00019049326,0.000013097773,0.000022151671,0.000021497837,0.000036659596,0.978973,0.0006570507,0.019218795,0.0008249923,0.000010481986],"about_ca_topic_score_codex":0.025207877,"about_ca_topic_score_gemma":0.03369194,"teacher_disagreement_score":0.025207877,"about_ca_system_score_codex":0.0010971889,"about_ca_system_score_gemma":0.0022472932,"threshold_uncertainty_score":0.05012232},"labels":[],"label_agreement":null},{"id":"W4407241185","doi":"10.1007/s10994-024-06670-4","title":"Schema-tune: noise-driven bias mitigation in transformer-based language models","year":2025,"lang":"en","type":"article","venue":"Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Schema (genetic algorithms); Transformer; Artificial intelligence; Natural language processing; Machine learning; Speech recognition; Engineering; Electrical engineering","score_opus":0.018024811983866115,"score_gpt":0.2643060661716281,"score_spread":0.24628125418776198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407241185","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01330603,0.00027331794,0.9765866,0.0002403531,0.00010129589,0.00006321753,0.00061740674,0.0077724005,0.0010394602],"genre_scores_gemma":[0.53861666,0.000510439,0.44924963,0.0005793054,0.00016431586,0.00023791453,0.0037938538,0.0028805032,0.0039673815],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99886984,0.0004668337,0.000070100665,0.00028876224,0.0002107476,0.000093648814],"domain_scores_gemma":[0.99724305,0.0014763693,0.00009693949,0.00066792505,0.00041289258,0.000102851605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023964972,0.0007791496,0.00083806727,0.00066150544,0.00036651196,0.001402978,0.001718941,0.0009971495,0.0044003073],"category_scores_gemma":[0.010637139,0.0005016929,0.0010022775,0.0007644337,0.00048512785,0.0026735628,0.0019372246,0.0021351185,0.0030341202],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012060427,0.00030915113,0.0043696202,0.0005877334,0.00043700848,0.0002750089,0.00073497294,0.16233422,0.040551923,0.054769386,0.033431347,0.70099366],"study_design_scores_gemma":[0.00006303724,0.000056581004,0.00024769016,0.000019532588,0.000059962116,0.000097024524,0.0000648436,0.9461794,0.013318113,0.03537262,0.004501036,0.000020207064],"about_ca_topic_score_codex":0.002324363,"about_ca_topic_score_gemma":0.0050523365,"teacher_disagreement_score":0.0044003073,"about_ca_system_score_codex":0.0005232056,"about_ca_system_score_gemma":0.0012847253,"threshold_uncertainty_score":0.0147204995},"labels":[],"label_agreement":null},{"id":"W4407309364","doi":"10.48550/arxiv.2502.04689","title":"Improving Language Models with Intentional Analysis","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Question answering; Computer science; Natural language processing; Language model; Artificial intelligence; Information retrieval","score_opus":0.033614337938902734,"score_gpt":0.26426984434759765,"score_spread":0.23065550640869492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407309364","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029875245,0.0004251095,0.95617324,0.0010099874,0.0000705728,0.00014236518,0.00053127034,0.008962345,0.0028098894],"genre_scores_gemma":[0.37857097,0.00041210494,0.6135015,0.0004199019,0.000081891616,0.00029009272,0.002760143,0.001232479,0.0027308692],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996985,0.001614391,0.00020088008,0.0005409332,0.0005150044,0.00014382343],"domain_scores_gemma":[0.98842245,0.008312497,0.0004671575,0.0017331435,0.0008343779,0.00023035836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041656825,0.0014696171,0.00080704316,0.0016354956,0.00062088046,0.0035941983,0.0019083332,0.0011344489,0.0041570277],"category_scores_gemma":[0.02251601,0.000633743,0.0021246083,0.0008929831,0.0009241553,0.0069148513,0.003851357,0.0027612874,0.0019409531],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044364622,0.0004354867,0.008441993,0.0010468763,0.0003637944,0.00023152732,0.0028283685,0.33617285,0.01585806,0.09032764,0.018548803,0.52530104],"study_design_scores_gemma":[0.000033452434,0.000043503434,0.00018715422,0.000037415106,0.000046620728,0.000025976668,0.00014752617,0.94167495,0.0024627966,0.05097254,0.0043482613,0.000019809066],"about_ca_topic_score_codex":0.005790345,"about_ca_topic_score_gemma":0.010328776,"teacher_disagreement_score":0.005790345,"about_ca_system_score_codex":0.0015086526,"about_ca_system_score_gemma":0.002682886,"threshold_uncertainty_score":0.022030473},"labels":[],"label_agreement":null},{"id":"W4407355838","doi":"10.1145/3709681","title":"Dialogue Benchmark Generation from Knowledge Graphs with Cost-Effective Retrieval-Augmented LLMs","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on Management of Data","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Benchmark (surveying); Computer science; Artificial intelligence; Information retrieval; Geography","score_opus":0.05128988014828679,"score_gpt":0.2954122021952202,"score_spread":0.2441223220469334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407355838","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0931278,0.0006280868,0.78108895,0.0005263693,0.00021434792,0.000959839,0.0027669207,0.114560336,0.0061272713],"genre_scores_gemma":[0.52473414,0.00016546823,0.45116344,0.0002994257,0.00006845826,0.0011054464,0.014242655,0.004989411,0.0032315583],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99592835,0.002090082,0.00024693293,0.0007815512,0.00073061796,0.00022245581],"domain_scores_gemma":[0.98882747,0.0067569925,0.00040366568,0.0020180522,0.0016901554,0.00030379437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003410154,0.0020921822,0.0012396389,0.0022651465,0.0007089255,0.0017398016,0.00282885,0.0011607552,0.00489074],"category_scores_gemma":[0.022113223,0.00065274374,0.00089632807,0.0011675162,0.00075791276,0.0027543015,0.0032713104,0.0013647428,0.0023988765],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018693177,0.00083767535,0.004778986,0.0011651808,0.00020971383,0.000816551,0.0014383157,0.23298718,0.039030798,0.009083797,0.044714082,0.66306835],"study_design_scores_gemma":[0.00016480048,0.00024966258,0.0007474704,0.00002982591,0.000059350416,0.000092291506,0.00031863645,0.9578103,0.022603178,0.009819552,0.008062644,0.000042190757],"about_ca_topic_score_codex":0.0045393924,"about_ca_topic_score_gemma":0.006106475,"teacher_disagreement_score":0.00489074,"about_ca_system_score_codex":0.0011495715,"about_ca_system_score_gemma":0.0014695527,"threshold_uncertainty_score":0.018034875},"labels":[],"label_agreement":null},{"id":"W4407423530","doi":"10.1016/j.enganabound.2025.106150","title":"Optimizing chatbot responsiveness: Automated history context selector via three-way decision for multi-turn dialogue Large Language Models","year":2025,"lang":"en","type":"article","venue":"Engineering Analysis with Boundary Elements","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Science and Technology Service Network Plan; Key Science and Technology Program of Shaanxi Province; National Key Research and Development Program of China; National Natural Science Foundation of China; Department of Science and Technology of Sichuan Province; Organization Department of Sichuan Provincial Party Committee; Ministry of Science and Technology of the People's Republic of China","keywords":"Chatbot; Computer science; Context (archaeology); Language model; Turn-taking; Natural language processing; Artificial intelligence; Linguistics; Philosophy; History; Conversation","score_opus":0.015600098102356824,"score_gpt":0.25577258355126214,"score_spread":0.24017248544890532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407423530","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06209543,0.0001680952,0.91629696,0.00029112658,0.00007594745,0.00016789397,0.00019866992,0.018060049,0.0026458004],"genre_scores_gemma":[0.75139743,0.000057120244,0.24218553,0.00015368054,0.00003581372,0.00021760115,0.0005441479,0.0016583772,0.0037502155],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99775153,0.000872282,0.000107643384,0.0005292079,0.00040902809,0.0003302969],"domain_scores_gemma":[0.9963677,0.0025272174,0.000107480664,0.00039405935,0.0003215964,0.00028190578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026513,0.0013238733,0.0017956507,0.0007532168,0.0011822601,0.0020360085,0.0022212637,0.0017480063,0.00799803],"category_scores_gemma":[0.007832302,0.0009229436,0.0012316481,0.00039966698,0.00087821274,0.002438777,0.003721703,0.0021140259,0.0022950002],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026986806,0.0007971843,0.004641495,0.00043124895,0.00025257186,0.0004974293,0.0015555789,0.4064423,0.049827144,0.018804807,0.014532212,0.4995194],"study_design_scores_gemma":[0.000020686986,0.000033033,0.000088135144,0.0000048086963,0.000014813183,0.000009290617,0.00007834598,0.992665,0.0033938938,0.003160128,0.0005214005,0.000010480685],"about_ca_topic_score_codex":0.0070573,"about_ca_topic_score_gemma":0.009940815,"teacher_disagreement_score":0.00799803,"about_ca_system_score_codex":0.0010990724,"about_ca_system_score_gemma":0.002163526,"threshold_uncertainty_score":0.026756048},"labels":[],"label_agreement":null},{"id":"W4407441485","doi":"10.1016/j.jslw.2025.101187","title":"Investigating L2 writers' critical AI literacy in AI-assisted writing: An APSE model","year":2025,"lang":"en","type":"article","venue":"Journal of Second Language Writing","topic":"Topic Modeling","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Literacy; Computer science; Natural language processing; Psychology; Pedagogy","score_opus":0.020744357858984364,"score_gpt":0.33690102792664295,"score_spread":0.3161566700676586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407441485","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.796872,0.0003020768,0.08596734,0.004799878,0.000035009587,0.0004966815,0.00013688322,0.00016619358,0.11122391],"genre_scores_gemma":[0.9873585,0.00009346605,0.010038466,0.00015558768,0.0000064878973,0.00020770232,0.000031291205,0.000016129774,0.0020922467],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.99470735,0.0031340383,0.00024236029,0.00063951535,0.0009480801,0.00032868455],"domain_scores_gemma":[0.9730118,0.018858645,0.0028443472,0.0017755958,0.0024112363,0.0010983514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059776637,0.0007505539,0.00032774988,0.0028397199,0.0017199167,0.0077938475,0.001434829,0.0022414941,0.0033745912],"category_scores_gemma":[0.023334611,0.00060904224,0.0006907439,0.0013083296,0.011323033,0.008108068,0.0056409216,0.0023041023,0.000643085],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022305662,0.0010037794,0.13140139,0.00096672826,0.00012247698,0.0019677838,0.3988118,0.0064134593,0.007721647,0.36277083,0.0013604102,0.08723666],"study_design_scores_gemma":[0.00018308351,0.00094295555,0.08303096,0.0013317767,0.00022927119,0.0029804208,0.39465156,0.09726683,0.011147098,0.36485466,0.043187857,0.00019347435],"about_ca_topic_score_codex":0.0023128495,"about_ca_topic_score_gemma":0.0019740334,"teacher_disagreement_score":0.0077938475,"about_ca_system_score_codex":0.0035719157,"about_ca_system_score_gemma":0.0040597618,"threshold_uncertainty_score":0.03161329},"labels":[],"label_agreement":null},{"id":"W4407570660","doi":"10.1145/3706598.3714020","title":"Fostering Appropriate Reliance on Large Language Models: The Role of Explanations, Sources, and Inconsistencies","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"Princeton University; National Science Foundation","keywords":"Key (lock); Computer science; Scale (ratio); Raising (metalworking); Cognitive psychology; Psychology; Computer security; Engineering","score_opus":0.03481069476889895,"score_gpt":0.2591592520451504,"score_spread":0.22434855727625144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407570660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8413632,0.00054006686,0.14820613,0.0018001405,0.000039207545,0.000251899,0.00009732321,0.0027230205,0.004978984],"genre_scores_gemma":[0.93330294,0.00015408898,0.06535483,0.00023973006,0.00002331542,0.00015337806,0.00008965082,0.00023809278,0.0004438806],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9732473,0.01995497,0.0014367942,0.0015334451,0.0033612587,0.00046631377],"domain_scores_gemma":[0.61952096,0.33103734,0.016485011,0.024276234,0.007118625,0.0015617969],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02733528,0.0009112439,0.0006807227,0.0010508302,0.00074346445,0.004714951,0.001531808,0.0017131397,0.0016503966],"category_scores_gemma":[0.21056257,0.0008434368,0.00045935958,0.0007121039,0.0015744001,0.005655042,0.0041227485,0.002363233,0.00047361996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017372125,0.0012126575,0.16254188,0.003493648,0.00038560014,0.0013759319,0.17127499,0.007636584,0.16006124,0.017050974,0.0036185163,0.46961075],"study_design_scores_gemma":[0.0010030447,0.0040611885,0.2290699,0.00442586,0.0017472378,0.0061075618,0.06731174,0.2593377,0.20581128,0.13545816,0.08433752,0.0013289088],"about_ca_topic_score_codex":0.0007896567,"about_ca_topic_score_gemma":0.0012798609,"teacher_disagreement_score":0.02733528,"about_ca_system_score_codex":0.0006956666,"about_ca_system_score_gemma":0.0016406564,"threshold_uncertainty_score":0.14456451},"labels":[],"label_agreement":null},{"id":"W4407588263","doi":"10.54254/2755-2721/2024.20851","title":"Comparative Analysis of Improved Versions of BERT Models on Chinese NLP Tasks","year":2025,"lang":"en","type":"article","venue":"Applied and Computational Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Earl Haig Secondary School","funders":"","keywords":"Natural language processing; Artificial intelligence; Computer science; Linguistics; Philosophy","score_opus":0.009348161555759269,"score_gpt":0.23727091577846823,"score_spread":0.22792275422270897,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407588263","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77287763,0.025795763,0.12450753,0.006926127,0.0022165717,0.0005886949,0.010608465,0.016939607,0.039539624],"genre_scores_gemma":[0.9098757,0.003496165,0.049287967,0.00076466677,0.00039157172,0.00038908026,0.022373414,0.00078338577,0.0126381125],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976708,0.0010215678,0.00016960497,0.00049062737,0.0003674724,0.0002798773],"domain_scores_gemma":[0.9888343,0.007693953,0.0002470293,0.0010738858,0.0017716617,0.0003790802],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073080757,0.0023784346,0.0015977748,0.00198963,0.0010002828,0.0014579293,0.002562924,0.0017275327,0.0035736896],"category_scores_gemma":[0.014274572,0.0005933905,0.0012992251,0.0023866042,0.00064171234,0.005222469,0.0013206814,0.0027200875,0.0017873512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023757638,0.00092635176,0.008417529,0.00091098127,0.00066217966,0.00022636654,0.0003950503,0.6209872,0.002264488,0.005856541,0.046170678,0.31080687],"study_design_scores_gemma":[0.00008223506,0.00029388012,0.002661954,0.000037720252,0.00013223596,0.000047052934,0.00009667267,0.99019027,0.0012431035,0.0022118478,0.0029542465,0.00004882467],"about_ca_topic_score_codex":0.056053285,"about_ca_topic_score_gemma":0.06370008,"teacher_disagreement_score":0.056053285,"about_ca_system_score_codex":0.003944762,"about_ca_system_score_gemma":0.002714845,"threshold_uncertainty_score":0.11145401},"labels":[],"label_agreement":null},{"id":"W4407681763","doi":"10.1145/3641554.3701844","title":"Integrating Small Language Models with Retrieval-Augmented Generation in Computing Education: Key Takeaways, Setup, and Practical Insights","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; University of Toronto","funders":"Universitas Brawijaya","keywords":"Key (lock); Computer science; Multimedia; Data science; Information retrieval; Operating system","score_opus":0.03359465857662799,"score_gpt":0.29085798431634574,"score_spread":0.2572633257397178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407681763","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2098492,0.0015977414,0.7459945,0.008622383,0.0003298637,0.00070602825,0.0007023843,0.020728234,0.011469761],"genre_scores_gemma":[0.52139145,0.0006417185,0.47057733,0.0009135448,0.00009226475,0.00035258738,0.0009928084,0.0011754987,0.003862755],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99721766,0.0012989229,0.00017978011,0.00042377133,0.000573493,0.0003064679],"domain_scores_gemma":[0.99580944,0.0018554386,0.000085739644,0.0015251377,0.00044381458,0.00028039046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004735054,0.00094499846,0.00064405904,0.0004659003,0.0006186795,0.0027176256,0.002270944,0.0011760192,0.0033968822],"category_scores_gemma":[0.0131750805,0.00059144536,0.0007037082,0.0004863138,0.0014220116,0.00689247,0.0036888812,0.0033293702,0.001762806],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00162934,0.0017763256,0.016614752,0.000731726,0.00020233798,0.0009552881,0.0040197005,0.14123166,0.05999972,0.065038405,0.027251368,0.6805494],"study_design_scores_gemma":[0.00022340208,0.00079627347,0.0015852153,0.000113893664,0.00006973357,0.00039302255,0.0011444581,0.86967623,0.042400252,0.042863354,0.04058744,0.0001467481],"about_ca_topic_score_codex":0.007013117,"about_ca_topic_score_gemma":0.009267921,"teacher_disagreement_score":0.007013117,"about_ca_system_score_codex":0.00099009,"about_ca_system_score_gemma":0.0020196203,"threshold_uncertainty_score":0.02504164},"labels":[],"label_agreement":null},{"id":"W4407737256","doi":"10.1109/ickg63256.2024.00017","title":"Axolotl: Fairness through Assisted Prompt Rewriting of Large Language Model Outputs","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Science Foundation","keywords":"Axolotl; Rewriting; Computer science; Programming language; Regeneration (biology)","score_opus":0.03888180439271441,"score_gpt":0.31185065349797914,"score_spread":0.27296884910526475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407737256","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019866059,0.00026029613,0.94936067,0.0007269527,0.0001562281,0.0001706336,0.00043443032,0.026386287,0.0026383626],"genre_scores_gemma":[0.5630055,0.00020159085,0.41755468,0.0012695747,0.00023998423,0.00043263665,0.0018539109,0.0062498376,0.009192204],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9922362,0.0033921155,0.00048942846,0.0015244291,0.0018595388,0.0004982441],"domain_scores_gemma":[0.9790253,0.011153448,0.0009648801,0.0066038934,0.001821108,0.0004314441],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009792542,0.0012162924,0.0010673134,0.0007728522,0.0011802595,0.0031284227,0.0037092944,0.0015481146,0.0072752996],"category_scores_gemma":[0.047660768,0.000639122,0.0012686834,0.00051028247,0.0020149567,0.004730717,0.006586726,0.0029700496,0.003611858],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023664325,0.0004094106,0.010748549,0.00081710384,0.00033022606,0.0010869453,0.0039573787,0.1965827,0.052131422,0.17104311,0.055822995,0.5047037],"study_design_scores_gemma":[0.000096982294,0.00013463631,0.0003284065,0.000045937497,0.000050150586,0.00014280144,0.00020125021,0.8622759,0.026224095,0.0952637,0.015180505,0.000055590575],"about_ca_topic_score_codex":0.0036318847,"about_ca_topic_score_gemma":0.005620359,"teacher_disagreement_score":0.009792542,"about_ca_system_score_codex":0.0013871177,"about_ca_system_score_gemma":0.0033671325,"threshold_uncertainty_score":0.05178851},"labels":[],"label_agreement":null},{"id":"W4407771189","doi":"10.1145/3641554.3701917","title":"Quantitative Evaluation of Using Large Language Models and Retrieval-Augmented Generation in Computer Science Education","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; University of British Columbia","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Language model; Information retrieval","score_opus":0.09692513595673274,"score_gpt":0.3935415129644818,"score_spread":0.29661637700774907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407771189","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9483604,0.0012994487,0.035953015,0.00059216726,0.00014469346,0.0012257879,0.0014051978,0.0032845559,0.007734646],"genre_scores_gemma":[0.94565314,0.00033586822,0.047853556,0.00015335548,0.000072069146,0.00072012533,0.0019817238,0.00035675502,0.0028735101],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9839252,0.011281198,0.0009940963,0.0011903507,0.002314067,0.00029513572],"domain_scores_gemma":[0.83461326,0.14184296,0.004818716,0.008911294,0.008050614,0.0017631383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015598651,0.00091139023,0.00059110887,0.0011624924,0.0005066898,0.001707999,0.0012516058,0.0012660519,0.0038308327],"category_scores_gemma":[0.09012094,0.00033170273,0.00051030534,0.0009190019,0.00089490286,0.002719821,0.0014830698,0.0011612057,0.0011316701],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.017798766,0.014321784,0.045941293,0.0060534296,0.0006258658,0.0002674887,0.005200346,0.079591475,0.057358526,0.0036549454,0.012951323,0.75623477],"study_design_scores_gemma":[0.0037521375,0.051998198,0.14766423,0.0006777379,0.0014820877,0.00080960424,0.00481001,0.6078638,0.13497327,0.0086826505,0.036751986,0.000534256],"about_ca_topic_score_codex":0.0023194607,"about_ca_topic_score_gemma":0.003084522,"teacher_disagreement_score":0.015598651,"about_ca_system_score_codex":0.0010834462,"about_ca_system_score_gemma":0.000790845,"threshold_uncertainty_score":0.08249456},"labels":[],"label_agreement":null},{"id":"W4407779872","doi":"10.1007/s42486-024-00173-w","title":"Abstractive summarization-based academic paper title drafting","year":2025,"lang":"en","type":"article","venue":"CCF Transactions on Pervasive Computing and Interaction","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval; Natural language processing","score_opus":0.016760335938860192,"score_gpt":0.2915291648003413,"score_spread":0.27476882886148113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407779872","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012382453,0.0033189373,0.79646844,0.008036833,0.013796134,0.0032300018,0.02607638,0.03264519,0.10404557],"genre_scores_gemma":[0.12073876,0.002977271,0.6329502,0.0015488996,0.005754272,0.0021370838,0.06660736,0.009104438,0.15818177],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941603,0.0020269854,0.00066800026,0.0011818318,0.0016617046,0.00030118623],"domain_scores_gemma":[0.9781917,0.006067254,0.00077785336,0.0031951333,0.011209696,0.0005583593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039154207,0.0015676268,0.0013464842,0.0050269943,0.002075735,0.005272761,0.0019264807,0.0014231585,0.103106126],"category_scores_gemma":[0.02258078,0.00079074636,0.001589517,0.0039099324,0.0008709602,0.0046907635,0.003258955,0.0019770532,0.062531605],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068520906,0.00015184075,0.0006106476,0.001756538,0.0001341853,0.0003077666,0.001410632,0.002973252,0.020153895,0.026139371,0.45800892,0.48766774],"study_design_scores_gemma":[0.00021410393,0.00033861966,0.0015211512,0.00059016456,0.0003150741,0.00037519512,0.0014839488,0.044026263,0.03675526,0.033900928,0.8803174,0.0001619216],"about_ca_topic_score_codex":0.0028966335,"about_ca_topic_score_gemma":0.0034855155,"teacher_disagreement_score":0.103106126,"about_ca_system_score_codex":0.0010231775,"about_ca_system_score_gemma":0.0028160026,"threshold_uncertainty_score":0.34492433},"labels":[],"label_agreement":null},{"id":"W4407953147","doi":"10.1145/3701551.3703483","title":"Bridging Historical Subgraph Optimization and Modern Graph Neural Network Approaches in Team Recommendation","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"Universitas Brawijaya","keywords":"Computer science; Bridging (networking); Graph; Artificial neural network; Subgraph isomorphism problem; Artificial intelligence; Theoretical computer science; Computer network","score_opus":0.0393111048748911,"score_gpt":0.23200218808673154,"score_spread":0.19269108321184045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407953147","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030741366,0.0027804137,0.9608259,0.001949134,0.00020089254,0.000046275924,0.00036423857,0.00025593714,0.0028358027],"genre_scores_gemma":[0.76808345,0.0044145905,0.21144147,0.0007206212,0.0014338584,0.00021314253,0.001807711,0.0003564059,0.011528773],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99822956,0.0009701564,0.00007578545,0.00037961628,0.00022357753,0.00012133415],"domain_scores_gemma":[0.9913522,0.006553509,0.0004130866,0.0007582205,0.0006652047,0.00025778046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034757738,0.0009895972,0.0019623383,0.0027341086,0.0009036246,0.0015792466,0.003665967,0.002183621,0.0041524167],"category_scores_gemma":[0.01486173,0.0010138921,0.0013235487,0.0039563654,0.0013095527,0.005349458,0.0014056072,0.003009758,0.0007225217],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015789465,0.00019889203,0.003431406,0.00027198644,0.00024724318,0.00006941977,0.00027920303,0.77999866,0.0004839568,0.102541655,0.008526849,0.10379279],"study_design_scores_gemma":[0.000007765104,0.000011884738,0.0002674149,0.000015245367,0.000020361642,0.0000095929945,0.000018212571,0.9472058,0.00008438684,0.05139354,0.00095947715,0.00000636999],"about_ca_topic_score_codex":0.015749738,"about_ca_topic_score_gemma":0.028333778,"teacher_disagreement_score":0.015749738,"about_ca_system_score_codex":0.0018309376,"about_ca_system_score_gemma":0.0010614212,"threshold_uncertainty_score":0.03131616},"labels":[],"label_agreement":null},{"id":"W4407953497","doi":"10.1145/3706628.3708827","title":"wa-hls4ml and lui-gnn: A Benchmark and GNN based Surrogate Model for hls4ml Resource and Latency Estimation","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"Universitas Brawijaya","keywords":"Benchmark (surveying); Latency (audio); Computer science; Surrogate model; Resource (disambiguation); Machine learning; Computer network; Geography; Telecommunications; Cartography","score_opus":0.014856595442261887,"score_gpt":0.2492664553552949,"score_spread":0.234409859913033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407953497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22681117,0.001455715,0.7233169,0.001175089,0.0003430192,0.00028176542,0.004695647,0.0059948596,0.035925854],"genre_scores_gemma":[0.81417817,0.0003511228,0.17240329,0.00021213487,0.000029079452,0.00029307778,0.0047803028,0.00047505237,0.007277811],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996055,0.00013689407,0.000020737967,0.00004843373,0.00013816686,0.00005039468],"domain_scores_gemma":[0.9988858,0.00057769153,0.00006994652,0.00014153456,0.00027561001,0.000049460592],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011122237,0.0008009719,0.0005470466,0.0006750787,0.0003382429,0.00081065804,0.0014055703,0.0009966443,0.004803373],"category_scores_gemma":[0.003554019,0.00028120226,0.0005984762,0.000606321,0.0004099606,0.00064246135,0.00057777774,0.0008730486,0.0006414775],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000076942364,0.000041528117,0.0008126582,0.000081634586,0.000016111362,0.000044383025,0.000014402671,0.9845154,0.0009988558,0.0025458497,0.0017011563,0.009151081],"study_design_scores_gemma":[0.0000075529065,0.000025818463,0.0001165815,0.0000071824356,0.0000023733523,0.000008233854,0.0000066406587,0.9974523,0.00065324357,0.00088918046,0.00082760246,0.0000033215506],"about_ca_topic_score_codex":0.012137275,"about_ca_topic_score_gemma":0.014414573,"teacher_disagreement_score":0.012137275,"about_ca_system_score_codex":0.00083390786,"about_ca_system_score_gemma":0.0014276358,"threshold_uncertainty_score":0.024133205},"labels":[],"label_agreement":null},{"id":"W4407953544","doi":"10.1145/3701551.3705706","title":"LLM4Eval@WSDM 2025: Large Language Model for Evaluation in Information Retrieval","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); University of Waterloo","funders":"Universitas Brawijaya","keywords":"Computer science; Information retrieval; Artificial intelligence","score_opus":0.026750719387370083,"score_gpt":0.32430639357556357,"score_spread":0.2975556741881935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407953544","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017082682,0.0038013859,0.84925586,0.009135071,0.0035113231,0.0036635334,0.01588195,0.07534406,0.022324111],"genre_scores_gemma":[0.13976398,0.0012828887,0.74636745,0.0038879365,0.0009919198,0.006493067,0.05845746,0.016184941,0.026570369],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9662855,0.025759349,0.001318355,0.002028647,0.0039175972,0.0006904278],"domain_scores_gemma":[0.95399404,0.02860343,0.00078831485,0.009260568,0.0048374655,0.0025160806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04383498,0.0025751132,0.0023758023,0.0026496628,0.0014466515,0.006603014,0.005416617,0.004763175,0.03240196],"category_scores_gemma":[0.0707552,0.0013540399,0.002961422,0.00166261,0.001710059,0.0080379965,0.007154211,0.006804872,0.013709432],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018775343,0.0011629671,0.0016090829,0.0017584829,0.00072093087,0.00032200155,0.00069068983,0.024452863,0.008119121,0.03898517,0.5406101,0.37969103],"study_design_scores_gemma":[0.0013744235,0.0014814315,0.0029819915,0.00048803326,0.00025189065,0.000382488,0.00042644172,0.53620297,0.023212511,0.101130895,0.33170837,0.00035850916],"about_ca_topic_score_codex":0.011014495,"about_ca_topic_score_gemma":0.013670453,"teacher_disagreement_score":0.04383498,"about_ca_system_score_codex":0.0042796386,"about_ca_system_score_gemma":0.0041541364,"threshold_uncertainty_score":0.23182434},"labels":[],"label_agreement":null},{"id":"W4408051557","doi":"10.3390/electronics14050973","title":"Efficient Ensemble of Deep Neural Networks for Multimodal Punctuation Restoration and the Spontaneous Informal Speech Dataset","year":2025,"lang":"en","type":"article","venue":"Electronics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Punctuation; Computer science; Artificial neural network; Deep neural networks; Speech recognition; Recurrent neural network; Artificial intelligence; Natural language processing","score_opus":0.008051341541456172,"score_gpt":0.23946076502313587,"score_spread":0.2314094234816797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408051557","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5336342,0.0047733802,0.15989149,0.0020668965,0.002001122,0.0010623022,0.20313874,0.073951915,0.019479979],"genre_scores_gemma":[0.3224247,0.0005447057,0.09888042,0.00046032344,0.00017769294,0.00095628016,0.5582886,0.0015031068,0.016764153],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896264,0.00028358737,0.000070322996,0.00031188698,0.00025942613,0.00011217881],"domain_scores_gemma":[0.99881566,0.0002763668,0.000056407836,0.00041704482,0.00032978694,0.000104704704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014183411,0.0024783234,0.00087853114,0.0015142622,0.0008304023,0.0007501942,0.002087589,0.0013581378,0.0042548496],"category_scores_gemma":[0.0036927091,0.00034559026,0.000976103,0.0009445209,0.00052831194,0.0011328036,0.0017814052,0.0020359561,0.004841192],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016374107,0.0012200867,0.013230036,0.00094764616,0.00061230623,0.0015802133,0.0004597696,0.06972762,0.02700114,0.0026154334,0.3306118,0.5503565],"study_design_scores_gemma":[0.0005019226,0.000864752,0.030578217,0.00018997563,0.0002940835,0.001526485,0.0009295306,0.75056666,0.049846433,0.006598504,0.1577831,0.0003203112],"about_ca_topic_score_codex":0.013210124,"about_ca_topic_score_gemma":0.03444658,"teacher_disagreement_score":0.013210124,"about_ca_system_score_codex":0.0008636613,"about_ca_system_score_gemma":0.0012119556,"threshold_uncertainty_score":0.026266456},"labels":[],"label_agreement":null},{"id":"W4408120198","doi":"10.1145/3704137.3704181","title":"Optimizing and Evaluating Enterprise Retrieval-Augmented Generation (RAG): A Content Design Perspective","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Perspective (graphical); Computer science; Content (measure theory); Information retrieval; Artificial intelligence; Mathematics","score_opus":0.1633340000320916,"score_gpt":0.3343451455154014,"score_spread":0.1710111454833098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408120198","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19769101,0.0038290005,0.7743899,0.0015032087,0.00015719146,0.0007879545,0.0007378163,0.0075801923,0.0133238025],"genre_scores_gemma":[0.68959504,0.0006883971,0.30423295,0.00023544458,0.00011886213,0.00016893909,0.0010985897,0.0006370902,0.0032247386],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943672,0.003652421,0.00024867154,0.00049849675,0.00086200715,0.00037122192],"domain_scores_gemma":[0.98742676,0.00926495,0.00039039587,0.0015314835,0.0011629214,0.00022333681],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006674171,0.0012827305,0.0016151244,0.0014423989,0.0005387831,0.003044235,0.0021390417,0.00213475,0.003114613],"category_scores_gemma":[0.022279942,0.0006474353,0.0009264263,0.001485675,0.00087324664,0.0030937747,0.0014873397,0.0010431634,0.0010004181],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012216688,0.00059213676,0.005011757,0.00058433856,0.0002712285,0.00017240924,0.00027726224,0.65341955,0.016432632,0.011454419,0.0051105856,0.30545205],"study_design_scores_gemma":[0.0000665239,0.00031133468,0.00068781513,0.000013839464,0.00012452847,0.000055547658,0.00008835571,0.9840329,0.0077883205,0.0052906633,0.0015201042,0.000020046362],"about_ca_topic_score_codex":0.006836995,"about_ca_topic_score_gemma":0.00819753,"teacher_disagreement_score":0.006836995,"about_ca_system_score_codex":0.0014366953,"about_ca_system_score_gemma":0.0020314741,"threshold_uncertainty_score":0.035296798},"labels":[],"label_agreement":null},{"id":"W4408145959","doi":"10.1109/icmla61862.2024.00187","title":"Evaluating the Efficacy of Large Language Models in Automating Academic Peer Reviews","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Natural language processing","score_opus":0.131467812164552,"score_gpt":0.4317497698631278,"score_spread":0.3002819576985758,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408145959","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7672227,0.003587864,0.19800562,0.0017062296,0.00049981556,0.0017646119,0.002083874,0.018678738,0.0064506023],"genre_scores_gemma":[0.83318985,0.00045660144,0.1619207,0.00030491376,0.00012687263,0.00046781788,0.0019550943,0.00033487633,0.0012431912],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9828689,0.012075487,0.0013015799,0.0019687607,0.0015250279,0.00026024072],"domain_scores_gemma":[0.80146784,0.1754998,0.006042127,0.0073848367,0.008373608,0.0012317918],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.02791481,0.0013759342,0.00090267666,0.0023890836,0.0006725134,0.0030153824,0.0018034874,0.0017489135,0.0011364688],"category_scores_gemma":[0.10889855,0.00061090605,0.0009035643,0.00137267,0.00055629906,0.0028006511,0.0014831291,0.0015886507,0.0013015486],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0050838236,0.0030443582,0.055603053,0.003315912,0.0014367583,0.00054345693,0.0029355017,0.3102587,0.021679975,0.0033746245,0.012032088,0.58069175],"study_design_scores_gemma":[0.00021514235,0.00091773237,0.00551572,0.0000812455,0.00016276962,0.00014034884,0.0003162698,0.979936,0.0077718394,0.0018670583,0.0030075908,0.00006823397],"about_ca_topic_score_codex":0.006699264,"about_ca_topic_score_gemma":0.009259839,"teacher_disagreement_score":0.9720852,"about_ca_system_score_codex":0.0015882492,"about_ca_system_score_gemma":0.0024265423,"threshold_uncertainty_score":0.14762938},"labels":[],"label_agreement":null},{"id":"W4408146816","doi":"10.1109/icairc64177.2024.10900179","title":"Comparative Analysis of Listwise Reranking with Large Language Models in Limited-Resource Language Contexts","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Resource (disambiguation); Natural language processing; Language model; Artificial intelligence","score_opus":0.02163557701518332,"score_gpt":0.2885934192222355,"score_spread":0.2669578422070522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408146816","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75161,0.030249204,0.17404597,0.0027639887,0.0010818274,0.0006328725,0.0059781894,0.01799663,0.01564125],"genre_scores_gemma":[0.8589241,0.0020792366,0.12300133,0.0003248449,0.00044731612,0.00027903653,0.010352245,0.0008135957,0.0037783356],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9939262,0.0039217058,0.00045716585,0.00073290203,0.0006836182,0.00027844205],"domain_scores_gemma":[0.9620545,0.030766176,0.00092257076,0.002942125,0.0026700357,0.0006446604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010265279,0.0024414256,0.0016934332,0.0037745363,0.0009435131,0.0017542586,0.0015460881,0.0013028571,0.0026746541],"category_scores_gemma":[0.0330643,0.00039956468,0.0010032391,0.002634484,0.0006533331,0.0042846263,0.0012120053,0.0016717138,0.0022903702],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039619952,0.0013039551,0.020490313,0.0021876588,0.0016165079,0.00043543646,0.0006344097,0.42876917,0.009472378,0.005225624,0.031956475,0.49394602],"study_design_scores_gemma":[0.00023576575,0.000998625,0.0036903366,0.0000613274,0.00025369684,0.00016483026,0.00028982124,0.98224556,0.0049474575,0.0036415092,0.0034012198,0.00006976005],"about_ca_topic_score_codex":0.013610315,"about_ca_topic_score_gemma":0.02896985,"teacher_disagreement_score":0.013610315,"about_ca_system_score_codex":0.0011030466,"about_ca_system_score_gemma":0.0017395184,"threshold_uncertainty_score":0.054288626},"labels":[],"label_agreement":null},{"id":"W4408156536","doi":"10.1007/978-3-031-82481-4_10","title":"Robust Infidelity: When Faithfulness Measures on Masked Language Models Are Misleading","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Programming language; Speech recognition","score_opus":0.04817885230213251,"score_gpt":0.24880635822914315,"score_spread":0.20062750592701065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408156536","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039871305,0.00086299854,0.9495091,0.0021844932,0.0002289968,0.00006406307,0.00026198334,0.0013976464,0.0056193755],"genre_scores_gemma":[0.8111328,0.0005886642,0.17901224,0.0016009741,0.00065547845,0.00014844218,0.00060925056,0.0012882751,0.0049639116],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9906268,0.004815026,0.00053874357,0.0016069653,0.0018887536,0.0005237017],"domain_scores_gemma":[0.8830209,0.093945585,0.004781796,0.013583301,0.0036247533,0.0010436622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015386992,0.0013495602,0.0017820328,0.002392458,0.001319242,0.0040091565,0.0026296885,0.0044623767,0.005611324],"category_scores_gemma":[0.16529614,0.0014077607,0.0009792924,0.00218396,0.0048956256,0.013147181,0.006737126,0.005253346,0.0012033512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014847642,0.00013608656,0.010201292,0.000622595,0.00032054906,0.0014762348,0.0019691915,0.105830714,0.009877083,0.59107935,0.015182998,0.26181906],"study_design_scores_gemma":[0.000027418797,0.00008704673,0.00081901375,0.00009893909,0.000049391074,0.0004275679,0.00016586775,0.306421,0.005112996,0.68450034,0.0022445258,0.000045875473],"about_ca_topic_score_codex":0.0014630856,"about_ca_topic_score_gemma":0.0009991731,"teacher_disagreement_score":0.015386992,"about_ca_system_score_codex":0.0014335697,"about_ca_system_score_gemma":0.0012374718,"threshold_uncertainty_score":0.08137518},"labels":[],"label_agreement":null},{"id":"W4408158553","doi":"10.1200/cci-24-00143","title":"Using a Longformer Large Language Model for Segmenting Unstructured Cancer Pathology Reports","year":2025,"lang":"en","type":"article","venue":"JCO Clinical Cancer Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of British Columbia","funders":"","keywords":"Pathology; Market segmentation; Computer science; Natural language processing; Medicine; Artificial intelligence; Business","score_opus":0.0939272117815288,"score_gpt":0.4464086519327034,"score_spread":0.3524814401511746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408158553","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068085685,0.0005682973,0.9123838,0.0009268952,0.00013618119,0.00034090082,0.0022905045,0.013250573,0.0020172037],"genre_scores_gemma":[0.5386092,0.00059861067,0.4370279,0.00076638936,0.0001205917,0.000686927,0.0102161765,0.0010402864,0.010933944],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992034,0.00024421312,0.00006797164,0.00030755557,0.00012097037,0.000055779492],"domain_scores_gemma":[0.9970265,0.0015226472,0.00024011485,0.0002913261,0.0008312158,0.00008813121],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023134053,0.001309282,0.00050780544,0.001372906,0.00047749642,0.0013067418,0.0019880203,0.0012554455,0.0043539586],"category_scores_gemma":[0.0045109103,0.00059643196,0.0013681946,0.00084332307,0.0006406759,0.0027584343,0.0011209949,0.0017716071,0.002887429],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012948296,0.00047533488,0.012929527,0.00050357165,0.00022840023,0.00084688765,0.0012189335,0.4638649,0.040328678,0.009293939,0.01914918,0.4498659],"study_design_scores_gemma":[0.00003384817,0.00012426273,0.0009827212,0.000020119202,0.000056028566,0.0001798104,0.00008245898,0.983365,0.008176924,0.00306155,0.0038901493,0.000027152595],"about_ca_topic_score_codex":0.01802088,"about_ca_topic_score_gemma":0.029634466,"teacher_disagreement_score":0.01802088,"about_ca_system_score_codex":0.0019795464,"about_ca_system_score_gemma":0.0025451933,"threshold_uncertainty_score":0.03583193},"labels":[],"label_agreement":null},{"id":"W4408184531","doi":"10.1145/3722449.3722461","title":"Report on the 1st Workshop on Large Language Model for Evaluation in Information Retrieval (LLM4Eval 2024) at SIGIR 2024","year":2024,"lang":"en","type":"article","venue":"ACM SIGIR Forum","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada); University of Waterloo","funders":"","keywords":"Information retrieval; Computer science; Natural language processing","score_opus":0.03542714885077752,"score_gpt":0.3152205595817855,"score_spread":0.279793410731008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408184531","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02447876,0.068527244,0.33957908,0.17495807,0.14056748,0.0064318776,0.059255783,0.017600248,0.16860154],"genre_scores_gemma":[0.06056798,0.017040081,0.18203326,0.024022823,0.025872037,0.0068441094,0.11622341,0.018556317,0.54884],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97543156,0.011956901,0.000888828,0.0025674733,0.0076551097,0.0015000798],"domain_scores_gemma":[0.9404712,0.019612074,0.0011577213,0.0065934374,0.023594605,0.008570987],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.056747783,0.002755925,0.00270539,0.003051185,0.002400398,0.010915673,0.003298965,0.0045833834,0.10599889],"category_scores_gemma":[0.05765055,0.0012282103,0.0024941273,0.002178418,0.0013810536,0.011321667,0.01080446,0.008560435,0.077981375],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043510328,0.00033511434,0.00058916875,0.00030476373,0.00007327588,0.00007604408,0.0003511653,0.0006996499,0.0016192859,0.0036421844,0.91157585,0.08029833],"study_design_scores_gemma":[0.0002711882,0.00033775,0.0025449877,0.00053599494,0.00012093327,0.00014358963,0.0004587061,0.004171002,0.0034841122,0.012682356,0.97511655,0.00013274538],"about_ca_topic_score_codex":0.010685548,"about_ca_topic_score_gemma":0.0137337865,"teacher_disagreement_score":0.9432522,"about_ca_system_score_codex":0.0034468344,"about_ca_system_score_gemma":0.0066636773,"threshold_uncertainty_score":0.35460162},"labels":[],"label_agreement":null},{"id":"W4408222036","doi":"10.2139/ssrn.5170458","title":"Large Language Models for Conceptual Modeling: Assessment and Application Potential","year":2025,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Computer science; Conceptual model; Natural language processing","score_opus":0.015988259062352912,"score_gpt":0.30250119961842,"score_spread":0.2865129405560671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408222036","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020676726,0.00166496,0.9694799,0.002539612,0.00010465515,0.00014149203,0.00085644575,0.0018374339,0.002698748],"genre_scores_gemma":[0.41092092,0.0028337503,0.5767316,0.0006372616,0.00058817916,0.00091681,0.003031502,0.0010327232,0.0033073202],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9900203,0.007920815,0.00026848374,0.0006690045,0.0009582495,0.00016318906],"domain_scores_gemma":[0.8481413,0.1383699,0.0017604465,0.006717986,0.0036084282,0.001401949],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016053773,0.0013421157,0.0020063433,0.0030284224,0.001359709,0.006657219,0.0034017314,0.0022166586,0.007302892],"category_scores_gemma":[0.079605915,0.001191613,0.0028470708,0.0037269122,0.0017421386,0.011366359,0.0027341389,0.0041106143,0.0018041297],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008123213,0.00057095074,0.0065183183,0.0008462028,0.00080977386,0.00019071512,0.0013609272,0.30388808,0.0015100251,0.5021384,0.014298721,0.16705559],"study_design_scores_gemma":[0.000043305252,0.00004315477,0.00030474982,0.000048125457,0.00007143964,0.000036265075,0.00010370216,0.7705713,0.00025434536,0.22535129,0.0031480282,0.000024243536],"about_ca_topic_score_codex":0.010041479,"about_ca_topic_score_gemma":0.0114677055,"teacher_disagreement_score":0.016053773,"about_ca_system_score_codex":0.003230993,"about_ca_system_score_gemma":0.003359721,"threshold_uncertainty_score":0.08490151},"labels":[],"label_agreement":null},{"id":"W4408281372","doi":"10.1109/tpami.2025.3550032","title":"Correlated Topic Modeling for Short Texts in Spherical Embedding Spaces","year":2025,"lang":"en","type":"article","venue":"IEEE Transactions on Pattern Analysis and Machine Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Hypersphere; Embedding; Knowledge graph; Graph; Word (group theory); Coherence (philosophical gambling strategy); Topic model; Information retrieval; Theoretical computer science; Mathematics","score_opus":0.026015574846276767,"score_gpt":0.30290152063852355,"score_spread":0.2768859457922468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408281372","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0459863,0.0012863132,0.950248,0.00035770572,0.00006695048,0.00007346448,0.00061687385,0.00040069318,0.000963666],"genre_scores_gemma":[0.838382,0.0030852684,0.1445479,0.00024347269,0.0005202527,0.00061819906,0.004161468,0.0003256189,0.008115904],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99880815,0.0004909083,0.000088821755,0.00034554076,0.0001659871,0.00010064689],"domain_scores_gemma":[0.99652714,0.0022800148,0.00048442435,0.00026922885,0.00035034853,0.00008885264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017209657,0.0012160063,0.001039017,0.0021070433,0.0004150973,0.0015424936,0.001252573,0.0012323017,0.0016604245],"category_scores_gemma":[0.007071519,0.0005438443,0.0014618302,0.0026814207,0.00086972537,0.0034462186,0.0012244759,0.00176151,0.0010502852],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037754024,0.00014886442,0.005830307,0.0004449143,0.00028579475,0.00044240154,0.0013907439,0.7214229,0.006465861,0.10426741,0.0058187777,0.1531044],"study_design_scores_gemma":[0.000005356373,0.000014185445,0.00038710944,0.000008769054,0.000010652674,0.000026518843,0.000034109173,0.9824011,0.00022838476,0.016178543,0.0006973704,0.000007944159],"about_ca_topic_score_codex":0.005689499,"about_ca_topic_score_gemma":0.006375032,"teacher_disagreement_score":0.005689499,"about_ca_system_score_codex":0.0008933972,"about_ca_system_score_gemma":0.00060113624,"threshold_uncertainty_score":0.011312783},"labels":[],"label_agreement":null},{"id":"W4408307058","doi":"10.2196/69663","title":"Identifying Adverse Events in Outpatients With Prostate Cancer Using Pharmaceutical Care Records in Community Pharmacies: Application of Named Entity Recognition","year":2025,"lang":"en","type":"article","venue":"JMIR Cancer","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Medicine; Pharmacy; Prostate cancer; Pharmaceutical care; Adverse effect; Medical prescription; Pharmacovigilance; Medical record; Family medicine; Pharmacology; Cancer; Internal medicine","score_opus":0.06694827125537767,"score_gpt":0.4024948304460421,"score_spread":0.33554655919066445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408307058","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79787904,0.0043867715,0.14235921,0.0016754302,0.00022849294,0.0023301423,0.03795705,0.0075711682,0.0056126737],"genre_scores_gemma":[0.7612319,0.0016634039,0.19873092,0.00039809194,0.00014577027,0.00050851656,0.03555006,0.00010279691,0.0016685103],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964888,0.00083356944,0.0009954689,0.0010178938,0.0005479326,0.000116326904],"domain_scores_gemma":[0.98553395,0.0068768878,0.003271734,0.0015905739,0.0023499148,0.00037687822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038436134,0.0006389277,0.00054743805,0.00564671,0.0004854431,0.0011480149,0.00074957573,0.00086817663,0.001187082],"category_scores_gemma":[0.014422726,0.00020153075,0.00092609844,0.00347492,0.00020871249,0.0016825325,0.0011952166,0.0005374513,0.00073584943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001231144,0.0005236669,0.39004964,0.0025325476,0.0005434556,0.0030053814,0.0021470159,0.004182232,0.02853474,0.00058456603,0.008124294,0.55854136],"study_design_scores_gemma":[0.0001788483,0.00095724623,0.75277895,0.0005334228,0.0014331661,0.0049308543,0.0032618192,0.12947538,0.06939936,0.002002411,0.034716688,0.00033187275],"about_ca_topic_score_codex":0.0056933737,"about_ca_topic_score_gemma":0.008106461,"teacher_disagreement_score":0.0056933737,"about_ca_system_score_codex":0.0005916762,"about_ca_system_score_gemma":0.0010779703,"threshold_uncertainty_score":0.02032727},"labels":[],"label_agreement":null},{"id":"W4408348562","doi":"10.1016/j.geomat.2025.100055","title":"Dynamic Named Entity Recognition model to distinguish authors’ positions relative to mentioned locations","year":2025,"lang":"en","type":"article","venue":"GEOMATICA","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"HORIZON EUROPE European Institute of Innovation and Technology; Österreichische Forschungsförderungsgesellschaft","keywords":"Computer science; Artificial intelligence; Natural language processing; Speech recognition","score_opus":0.023290107782961134,"score_gpt":0.2971115254029344,"score_spread":0.27382141761997325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408348562","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13405383,0.0010955048,0.8143762,0.0023212216,0.0007088034,0.0005584411,0.019878633,0.0163398,0.010667665],"genre_scores_gemma":[0.62286794,0.00054295827,0.32961878,0.00064131524,0.00020516106,0.00057027896,0.028413737,0.0003726864,0.016767116],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99855036,0.0002717278,0.00017899716,0.00067905226,0.00021156207,0.00010828786],"domain_scores_gemma":[0.9964126,0.0018097456,0.0004154911,0.00041824573,0.0008410169,0.0001029382],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019234079,0.0010465817,0.00070454495,0.0026505487,0.0006615502,0.0013524869,0.0022294607,0.0014460235,0.0026222875],"category_scores_gemma":[0.0061719804,0.000417036,0.0014241298,0.0023640909,0.0004563322,0.003915294,0.0011346596,0.0018819487,0.004127299],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010408122,0.00058441074,0.051607314,0.0006593347,0.00033335225,0.0022567285,0.0017321241,0.2688063,0.020690123,0.025340376,0.057943758,0.56900537],"study_design_scores_gemma":[0.00001520688,0.0000432558,0.0034144584,0.000033650344,0.00006485234,0.00025671595,0.00012274297,0.97194797,0.005967781,0.006638318,0.011457965,0.00003715118],"about_ca_topic_score_codex":0.010971409,"about_ca_topic_score_gemma":0.017367356,"teacher_disagreement_score":0.010971409,"about_ca_system_score_codex":0.0011922619,"about_ca_system_score_gemma":0.0011688583,"threshold_uncertainty_score":0.021815062},"labels":[],"label_agreement":null},{"id":"W4408374252","doi":"10.1186/s41239-025-00513-5","title":"Correction: “Scarlet Cloak and the Forest Adventure”: a preliminary study of the impact of AI on commonly used writing tools","year":2025,"lang":"en","type":"article","venue":"International Journal of Educational Technology in Higher Education","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Adventure; Cloak; Computer science; Artificial intelligence; Optics; Physics; Metamaterial","score_opus":0.02034335034577278,"score_gpt":0.348758313584627,"score_spread":0.3284149632388542,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408374252","genre_codex":"editorial","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005332283,0.0005411793,0.00056565483,0.17055933,0.82547045,0.000028691795,0.000497575,0.00029025285,0.0015137043],"genre_scores_gemma":[0.0471366,0.0038116297,0.0041584363,0.30511758,0.49706322,0.00042144966,0.0012811898,0.0015209609,0.13948897],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.990168,0.0019998297,0.0017911755,0.001447576,0.003932104,0.00066138885],"domain_scores_gemma":[0.87886995,0.028330062,0.0068564275,0.005150962,0.07708363,0.0037089821],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006278417,0.0016938976,0.0015823797,0.00336442,0.0050905943,0.005061043,0.004132041,0.009616258,0.030149188],"category_scores_gemma":[0.16797376,0.0008578107,0.0012413361,0.0025595666,0.00466271,0.003158432,0.0030758684,0.016936986,0.017570503],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000032396812,0.0000064605556,0.00012764147,0.00011559346,0.000011507776,0.00033308018,0.00027683008,0.000019672854,0.000042141655,0.00046413948,0.9950907,0.0034798083],"study_design_scores_gemma":[0.0000706003,0.00003144393,0.0025546737,0.00072167575,0.000052993142,0.0012895357,0.0013393859,0.0003382868,0.00048913306,0.0010795929,0.99195147,0.00008118484],"about_ca_topic_score_codex":0.028886912,"about_ca_topic_score_gemma":0.03152423,"teacher_disagreement_score":0.030149188,"about_ca_system_score_codex":0.0056354543,"about_ca_system_score_gemma":0.007103787,"threshold_uncertainty_score":0.100859106},"labels":[],"label_agreement":null},{"id":"W4408521274","doi":"10.1109/bhi62660.2024.10913641","title":"Improving Interpretability of Radiology Report-based Pediatric Brain Tumor Pathology Classification and Key-phrases Extraction Using Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"SickKids Foundation; Hospital for Sick Children; University of Toronto","funders":"","keywords":"Interpretability; Computer science; Key (lock); Natural language processing; Artificial intelligence; Information extraction; Feature extraction; Pathology; Medicine","score_opus":0.028976896577114573,"score_gpt":0.30281162121157673,"score_spread":0.2738347246344622,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408521274","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5733164,0.0033818763,0.39566103,0.00221703,0.0003853789,0.00058275525,0.010370522,0.01061235,0.0034726036],"genre_scores_gemma":[0.8685199,0.00075189036,0.11455586,0.0002146393,0.00025836544,0.00028135692,0.014259864,0.00033120284,0.0008268913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951497,0.0021522644,0.0007264983,0.0009447318,0.0008078361,0.0002190374],"domain_scores_gemma":[0.9732851,0.020173924,0.0024372162,0.0016193137,0.0021924449,0.00029209757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077300137,0.0018436094,0.0007051791,0.0058048265,0.0004103526,0.002398534,0.0010365221,0.0010079728,0.001243256],"category_scores_gemma":[0.030827893,0.00033776247,0.0015103811,0.0020749609,0.0004346788,0.0023923852,0.0015932263,0.0013211123,0.0013117259],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014085372,0.00070176215,0.2260917,0.0012064287,0.0009098052,0.0021274486,0.0018242087,0.06763332,0.038064368,0.0021182955,0.019484296,0.63842976],"study_design_scores_gemma":[0.00009829872,0.00035507124,0.055574227,0.00022029768,0.0005401436,0.0014025619,0.0009152073,0.89934105,0.024972778,0.005941191,0.010523865,0.00011530401],"about_ca_topic_score_codex":0.0038864403,"about_ca_topic_score_gemma":0.003866178,"teacher_disagreement_score":0.0077300137,"about_ca_system_score_codex":0.0006925783,"about_ca_system_score_gemma":0.001297861,"threshold_uncertainty_score":0.04088068},"labels":[],"label_agreement":null},{"id":"W4408595711","doi":"10.1016/j.tics.2025.02.003","title":"Studying memory narratives with natural language processing","year":2025,"lang":"en","type":"review","venue":"Trends in Cognitive Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Baycrest Hospital; McGill University","funders":"Canada Excellence Research Chairs, Government of Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Psychology; Narrative; Cognitive science; Linguistics; Cognitive psychology; Natural (archaeology); Communication; History","score_opus":0.18630717812774528,"score_gpt":0.4533976597653287,"score_spread":0.2670904816375834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408595711","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00047807224,0.9969405,0.0009303842,0.0006276728,0.00013697214,0.0000099495,0.000027142336,0.000013075676,0.00083626725],"genre_scores_gemma":[0.00644534,0.9901398,0.0019191006,0.0004091794,0.0005390866,0.000035542715,0.00008812634,0.000008098422,0.00041579668],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99946195,0.0002009958,0.00005764533,0.00013271044,0.0001201252,0.000026625827],"domain_scores_gemma":[0.989892,0.008918216,0.0004558977,0.00012716113,0.0004908264,0.00011595855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019014458,0.0011200256,0.0015133108,0.0038736465,0.00028657308,0.0022718604,0.0011461388,0.0015595744,0.0029083814],"category_scores_gemma":[0.011283626,0.0003540721,0.00070680963,0.0038753415,0.0009930066,0.004350599,0.0010033733,0.001321807,0.0007879062],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005476833,0.00005511099,0.00095599133,0.01994191,0.00017068886,0.000055971228,0.0003378373,0.0003606556,0.00030344364,0.006515777,0.007935639,0.96331215],"study_design_scores_gemma":[0.00012494468,0.00023301993,0.013269805,0.058585875,0.0017750358,0.0014861022,0.0018928928,0.0020737727,0.0019107042,0.07050962,0.8480038,0.00013442014],"about_ca_topic_score_codex":0.0032982451,"about_ca_topic_score_gemma":0.0049908604,"teacher_disagreement_score":0.0038736465,"about_ca_system_score_codex":0.0011580212,"about_ca_system_score_gemma":0.0032622246,"threshold_uncertainty_score":0.010055959},"labels":[],"label_agreement":null},{"id":"W4408612569","doi":"10.2139/ssrn.5184646","title":"Leveraging Open-Source Large Language Models (Llms) to Enhance Data Extraction in Scoping Reviews: A Case Study on Disability and Ai Applications","year":2025,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Open source; Computer science; Data source; Data science; Data extraction; Natural language processing; World Wide Web; Information retrieval; Political science; MEDLINE; Programming language","score_opus":0.06356386152874574,"score_gpt":0.41210621814550086,"score_spread":0.34854235661675514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408612569","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.044646554,0.016720442,0.851549,0.0147578465,0.00065086136,0.00673887,0.03869855,0.015195565,0.011042211],"genre_scores_gemma":[0.12382618,0.0050905533,0.8411294,0.0011800642,0.00018362107,0.003603847,0.02155143,0.0015067075,0.0019282029],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9379632,0.043056075,0.010753973,0.003340547,0.004469229,0.00041706391],"domain_scores_gemma":[0.48437026,0.4593636,0.0170668,0.020717619,0.017215136,0.0012666298],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07584662,0.0015180323,0.0014854897,0.0159493,0.0017173928,0.007840084,0.0027603987,0.0020762742,0.004130813],"category_scores_gemma":[0.2855896,0.0010786641,0.0038870769,0.01373916,0.000882341,0.008197162,0.0062144063,0.002802026,0.0026262715],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091228646,0.00042486945,0.017625641,0.05592512,0.0021903627,0.001915876,0.025441365,0.011385141,0.011900831,0.022147004,0.04308958,0.8070419],"study_design_scores_gemma":[0.0006854617,0.00083444035,0.02156091,0.037552666,0.008478016,0.0024635368,0.017774576,0.1394438,0.041013822,0.17746885,0.5518227,0.0009011264],"about_ca_topic_score_codex":0.0055452953,"about_ca_topic_score_gemma":0.023559585,"teacher_disagreement_score":0.9241534,"about_ca_system_score_codex":0.002545128,"about_ca_system_score_gemma":0.014646633,"threshold_uncertainty_score":0.40112007},"labels":[],"label_agreement":null},{"id":"W4408688975","doi":"10.32388/b69sky","title":"DAFE: LLM-Based Evaluation Through Dynamic Arbitration for Free-Form Question-Answering","year":2025,"lang":"en","type":"preprint","venue":"Qeios","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Arbitration; Question answering; Computer science; Information retrieval; Political science; Law","score_opus":0.03800586777782443,"score_gpt":0.35895999524778194,"score_spread":0.3209541274699575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408688975","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016052714,0.00040996278,0.96725637,0.00049916765,0.00013989392,0.0006632912,0.0007760727,0.010801514,0.0034009675],"genre_scores_gemma":[0.3487172,0.0001625563,0.641537,0.00058227277,0.00012818439,0.0018714264,0.002572527,0.0012389697,0.0031899041],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96691763,0.021878067,0.0019190528,0.0036488073,0.004986011,0.0006503663],"domain_scores_gemma":[0.9517157,0.03482217,0.001950202,0.0049268263,0.00551625,0.0010688885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026862625,0.0022686115,0.0015263137,0.00363922,0.0009698757,0.0041174307,0.0032916605,0.0030292356,0.010311916],"category_scores_gemma":[0.10560308,0.00059968384,0.0011821132,0.0012479139,0.0016779725,0.0062692994,0.007627098,0.0028299126,0.0038724965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014981381,0.00077065604,0.008476933,0.0017288048,0.00044109978,0.00032854208,0.0045483867,0.054636937,0.027869347,0.046996787,0.034473926,0.81823045],"study_design_scores_gemma":[0.00023154858,0.00038080316,0.0025832949,0.00019224705,0.00007290838,0.00019626215,0.0006551268,0.8548994,0.01915394,0.098181985,0.023287201,0.00016525385],"about_ca_topic_score_codex":0.0019975386,"about_ca_topic_score_gemma":0.003320475,"teacher_disagreement_score":0.026862625,"about_ca_system_score_codex":0.0016878246,"about_ca_system_score_gemma":0.0018985547,"threshold_uncertainty_score":0.14206481},"labels":[],"label_agreement":null},{"id":"W4408781649","doi":"10.36227/techrxiv.174285217.79890225/v1","title":"Improving Large Language Model Performance Through Compression and Optimization","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Compression (physics); Materials science; Composite material","score_opus":0.019473796125071708,"score_gpt":0.26902898415700505,"score_spread":0.24955518803193333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408781649","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01544878,0.0050120233,0.96356523,0.0017565119,0.0001966055,0.00013002341,0.0012453053,0.007979453,0.004666091],"genre_scores_gemma":[0.27245885,0.01013202,0.7025707,0.0011793708,0.0003692077,0.00050235767,0.0065790405,0.0021972675,0.0040112245],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99741435,0.0008294542,0.0002208875,0.00039793472,0.00095265044,0.0001846246],"domain_scores_gemma":[0.9934691,0.0043593054,0.00019592968,0.0012814713,0.0005942208,0.00010003693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029792802,0.001617335,0.0015927356,0.0016533203,0.00067643955,0.0032643294,0.0031163634,0.001298157,0.0042368528],"category_scores_gemma":[0.020757923,0.0007624385,0.0015315004,0.0026945497,0.0010769607,0.007712328,0.002518065,0.0031188417,0.0028684882],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031190302,0.00020157744,0.002904663,0.0011795617,0.00025905133,0.00025521126,0.00045091973,0.36105394,0.009479772,0.060634643,0.025737854,0.53753096],"study_design_scores_gemma":[0.00003527047,0.00004931284,0.00024148295,0.000071265094,0.00005104567,0.000120346995,0.00009544522,0.94412464,0.004731365,0.04254394,0.007911358,0.000024474472],"about_ca_topic_score_codex":0.008036251,"about_ca_topic_score_gemma":0.013326056,"teacher_disagreement_score":0.008036251,"about_ca_system_score_codex":0.0015131516,"about_ca_system_score_gemma":0.002857996,"threshold_uncertainty_score":0.015978932},"labels":[],"label_agreement":null},{"id":"W4408834179","doi":"10.1051/itmconf/20257605013","title":"Natural Language Processing Techniques for Information Retrieval Enhancing Search Engines with Semantic Understanding","year":2025,"lang":"en","type":"article","venue":"ITM Web of Conferences","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Information retrieval; Search engine; Semantic search; Natural language processing; Human–computer information retrieval; Artificial intelligence","score_opus":0.026924476117215126,"score_gpt":0.2920163499823629,"score_spread":0.2650918738651478,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408834179","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046238312,0.0014472029,0.98726684,0.0011284001,0.00007116089,0.00015102261,0.00017058811,0.0009246555,0.0042164032],"genre_scores_gemma":[0.097653925,0.0018196483,0.89530456,0.00047854715,0.00020168338,0.00030242873,0.00048435756,0.00020735394,0.003547465],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99793327,0.0009169885,0.00018003595,0.0002539011,0.00063131156,0.00008457031],"domain_scores_gemma":[0.99708456,0.0017130354,0.00022789557,0.0005194891,0.00041239805,0.000042618485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035535134,0.00073952787,0.0008326141,0.0035107925,0.00070411066,0.0025597129,0.0012762218,0.0011507159,0.0037684808],"category_scores_gemma":[0.009613292,0.00055338483,0.0017671164,0.0028030388,0.0014031536,0.007414951,0.0018860607,0.0017087026,0.0020395876],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015003224,0.00026169993,0.001146301,0.0015501275,0.00018217787,0.0002400162,0.0014585088,0.044379998,0.023365932,0.46459925,0.015708731,0.44695723],"study_design_scores_gemma":[0.000067930145,0.00011817009,0.0006032868,0.0001486864,0.00012118199,0.00045004647,0.00044260654,0.4237481,0.012027852,0.5035173,0.058677435,0.000077278935],"about_ca_topic_score_codex":0.0016879805,"about_ca_topic_score_gemma":0.002588207,"teacher_disagreement_score":0.0037684808,"about_ca_system_score_codex":0.001069218,"about_ca_system_score_gemma":0.0010558169,"threshold_uncertainty_score":0.018792987},"labels":[],"label_agreement":null},{"id":"W4408846465","doi":"10.1145/3689031.3717461","title":"Mist: Efficient Distributed Training of Large Language Models via Memory-Parallelism Co-Optimization","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Mist; Parallel computing; Parallelism (grammar); Training (meteorology); Computer architecture; Distributed computing","score_opus":0.0171281055888019,"score_gpt":0.2690719334659288,"score_spread":0.2519438278771269,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408846465","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06532726,0.0007915922,0.8874012,0.00043297355,0.00013908352,0.00013531868,0.00039134035,0.039970074,0.0054112645],"genre_scores_gemma":[0.5543312,0.00025281616,0.43439686,0.00046478552,0.00008195491,0.0003986548,0.0020450682,0.0019352604,0.006093439],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994097,0.00012317472,0.000029569068,0.00018832105,0.00015435377,0.00009488062],"domain_scores_gemma":[0.9989611,0.00051486655,0.00006353834,0.00023038402,0.00016444607,0.0000655687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00059437397,0.0013204362,0.0008217448,0.00045590458,0.00058420375,0.00075611577,0.0022784425,0.00080092764,0.0038001684],"category_scores_gemma":[0.0023760092,0.0005244996,0.0009150164,0.0006966737,0.0006524564,0.0020234298,0.0015992983,0.0017040597,0.00148245],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005213397,0.00041171594,0.0023004038,0.00026966404,0.0001411597,0.0002539632,0.00025613257,0.5888264,0.022853643,0.0064650928,0.02172174,0.35597882],"study_design_scores_gemma":[0.000031949323,0.00004278971,0.00010680368,0.0000031587208,0.000008834072,0.000020626774,0.000020416504,0.99300635,0.0033845203,0.0023654997,0.0010036379,0.0000054200086],"about_ca_topic_score_codex":0.0077760317,"about_ca_topic_score_gemma":0.015384009,"teacher_disagreement_score":0.0077760317,"about_ca_system_score_codex":0.0007876554,"about_ca_system_score_gemma":0.0019568382,"threshold_uncertainty_score":0.015461504},"labels":[],"label_agreement":null},{"id":"W4408881088","doi":"10.1007/979-8-8688-1221-7_4","title":"Using AI in the Enterprise! Creating a Text Summarizer for Slack Messages","year":2025,"lang":"en","type":"book-chapter","venue":"Apress eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brandon University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence","score_opus":0.05639990912126809,"score_gpt":0.30068673264188317,"score_spread":0.2442868235206151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408881088","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010689409,0.015292188,0.671349,0.008407269,0.004948201,0.00024973176,0.0024076747,0.033346105,0.25331044],"genre_scores_gemma":[0.081408,0.008870102,0.41976887,0.0030370643,0.0032500965,0.00032894316,0.0057801893,0.0058903443,0.47166643],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99975103,0.000069895505,0.000014964012,0.0000592782,0.00008956718,0.00001517304],"domain_scores_gemma":[0.9991529,0.00052303355,0.000040475814,0.000082425264,0.00014237748,0.00005875228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004995386,0.0006711639,0.0003133445,0.0012001261,0.00059443945,0.0029331145,0.00057424075,0.0005957658,0.029396808],"category_scores_gemma":[0.0020912783,0.00025482147,0.00041075377,0.0011976451,0.0003233409,0.00399858,0.00085075415,0.0014067877,0.018921206],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006318507,0.00005910428,0.00021705407,0.0005431299,0.000034108685,0.00013932497,0.0019495292,0.0016607489,0.008097441,0.036350545,0.29152972,0.659356],"study_design_scores_gemma":[0.00001527468,0.00006345874,0.00044647543,0.00015539605,0.000033483007,0.00027028014,0.00070049294,0.012057649,0.0040026805,0.024938349,0.9572852,0.000031343523],"about_ca_topic_score_codex":0.0007281045,"about_ca_topic_score_gemma":0.0015311999,"teacher_disagreement_score":0.029396808,"about_ca_system_score_codex":0.00046985666,"about_ca_system_score_gemma":0.0003432334,"threshold_uncertainty_score":0.09834212},"labels":[],"label_agreement":null},{"id":"W4408912911","doi":"10.1016/j.datak.2025.102440","title":"Customized long short-term memory architecture for multi-document summarization with improved text feature set","year":2025,"lang":"en","type":"article","venue":"Data & Knowledge Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Automatic summarization; Computer science; Term (time); Feature (linguistics); Set (abstract data type); Information retrieval; Architecture; Natural language processing; Artificial intelligence; Programming language; History; Linguistics","score_opus":0.023473699877713024,"score_gpt":0.2815724607136532,"score_spread":0.2580987608359402,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408912911","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07622569,0.003327904,0.8527089,0.0003597753,0.00069491565,0.0003199404,0.0039363685,0.058898322,0.0035283226],"genre_scores_gemma":[0.38012818,0.0008823655,0.5864038,0.000524704,0.00043704634,0.00074454985,0.013350858,0.0008394371,0.016689075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99971324,0.00002908471,0.00004603561,0.000094105315,0.00007271098,0.0000448117],"domain_scores_gemma":[0.99932754,0.00014338367,0.000050497794,0.00013440034,0.00029810626,0.000046153087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041364605,0.0009054534,0.00085620984,0.0017896343,0.00058967876,0.0010729632,0.0016668164,0.0006282655,0.0087756235],"category_scores_gemma":[0.0010196218,0.00030578446,0.00059703144,0.002269811,0.00015159238,0.0014936435,0.0007630839,0.0006535944,0.0043267175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011295596,0.0003265448,0.0010697739,0.00030370487,0.0001864673,0.00017631822,0.00015546278,0.008972655,0.08531071,0.0011286136,0.025212532,0.8760276],"study_design_scores_gemma":[0.00028481727,0.0013084314,0.0047110203,0.00006506864,0.0005268432,0.00041754145,0.0002599387,0.76520306,0.18772696,0.0046844687,0.034664642,0.00014715375],"about_ca_topic_score_codex":0.0056710136,"about_ca_topic_score_gemma":0.009160771,"teacher_disagreement_score":0.0087756235,"about_ca_system_score_codex":0.0005435116,"about_ca_system_score_gemma":0.0010327755,"threshold_uncertainty_score":0.029357374},"labels":[],"label_agreement":null},{"id":"W4408928818","doi":"10.3389/fdgth.2025.1484521","title":"Interactive Panel Summaries of the 2024 Voice AI Symposium","year":2025,"lang":"en","type":"review","venue":"Frontiers in Digital Health","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Simon Fraser University","funders":"","keywords":"Panel discussion; Viewpoints; Event (particle physics); Conversation; Computer science; Transcription (linguistics); Multimedia; Dialog box; World Wide Web; Psychology; Visual arts; Communication; Linguistics","score_opus":0.02909328690872451,"score_gpt":0.3169712485560776,"score_spread":0.28787796164735313,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408928818","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035275568,0.5186202,0.009811201,0.06405136,0.24122998,0.0049698297,0.015779866,0.0010463938,0.14096363],"genre_scores_gemma":[0.023205552,0.45490167,0.014744125,0.044233654,0.13189599,0.010571623,0.026434813,0.0005332406,0.29347935],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981793,0.0006829153,0.00021088087,0.00014997281,0.0006285673,0.00014841964],"domain_scores_gemma":[0.99388784,0.0018349655,0.00059641333,0.00018429547,0.0026577187,0.00083873473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006480599,0.0011700161,0.0008644184,0.002716034,0.00072612654,0.0017818783,0.0013343621,0.0020098772,0.067779124],"category_scores_gemma":[0.009549003,0.00039362043,0.001274895,0.0022247827,0.00023453218,0.001413478,0.0022258805,0.0021344908,0.03340926],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016135041,0.000055069115,0.00004501637,0.0038865744,0.00003984105,0.00007808455,0.00014717063,0.00010486663,0.0008685854,0.00078187423,0.8888805,0.10495105],"study_design_scores_gemma":[0.000021515574,0.000045834713,0.00032387365,0.0009290692,0.00002054235,0.000055985303,0.000056501838,0.000019693365,0.00016161338,0.00032406274,0.99803287,0.000008464326],"about_ca_topic_score_codex":0.0007062804,"about_ca_topic_score_gemma":0.0024451856,"teacher_disagreement_score":0.067779124,"about_ca_system_score_codex":0.0010095249,"about_ca_system_score_gemma":0.002382466,"threshold_uncertainty_score":0.22674376},"labels":[],"label_agreement":null},{"id":"W4409009926","doi":"10.36227/techrxiv.174345061.15598909/v1","title":"The Next Frontier in AI Research with Distributed and Multimodal Large Language Models","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Frontier; Computer science; Natural language processing; Artificial intelligence; History; Archaeology","score_opus":0.05766419187475675,"score_gpt":0.34150658655928434,"score_spread":0.2838423946845276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409009926","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0047634803,0.019136792,0.9431935,0.020503266,0.00036769637,0.00005252691,0.00023942518,0.001184736,0.0105586415],"genre_scores_gemma":[0.2242294,0.04896438,0.7096831,0.0036681045,0.0020403287,0.00041714162,0.0012195304,0.0010151787,0.0087628225],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99790996,0.0010386066,0.00009029372,0.0003490609,0.00050794694,0.00010410846],"domain_scores_gemma":[0.9933257,0.004061294,0.00017367493,0.0016900364,0.0005373658,0.00021186113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003985555,0.00076333195,0.0011715859,0.0008976741,0.0006536535,0.004303258,0.002323282,0.001776083,0.005453191],"category_scores_gemma":[0.009718622,0.0005985461,0.001006636,0.0021271291,0.0023153564,0.011767829,0.002829517,0.004677456,0.0023135],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000088423876,0.00009914098,0.0011425317,0.0011352174,0.00018221137,0.000109664164,0.00056327926,0.0804811,0.0024770692,0.55368364,0.030572075,0.3294657],"study_design_scores_gemma":[0.0000160517,0.000038192014,0.00028703103,0.00014258383,0.00002690213,0.00008965683,0.00022313808,0.19792078,0.0008847087,0.7309405,0.069404595,0.000025868245],"about_ca_topic_score_codex":0.0026355442,"about_ca_topic_score_gemma":0.0032253803,"teacher_disagreement_score":0.005453191,"about_ca_system_score_codex":0.0018707874,"about_ca_system_score_gemma":0.002226874,"threshold_uncertainty_score":0.021077871},"labels":[],"label_agreement":null},{"id":"W4409047962","doi":"10.1145/3721146.3721940","title":"Accelerating MoE Model Inference with Expert Sharding","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Inference; Computer science; Artificial intelligence","score_opus":0.06416343562599425,"score_gpt":0.3062068170236703,"score_spread":0.24204338139767606,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409047962","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036641873,0.00026131142,0.9427996,0.00029875024,0.00009821623,0.000054671928,0.00034002046,0.017342482,0.0021630933],"genre_scores_gemma":[0.46614218,0.00016600505,0.5244902,0.0004204661,0.000077214325,0.00010681653,0.0014998573,0.0013138864,0.005783461],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993412,0.00012774748,0.000041840685,0.00020726265,0.00017394527,0.00010799044],"domain_scores_gemma":[0.9983432,0.00067918666,0.00009092583,0.000525909,0.00023864828,0.00012217964],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00092959433,0.0009936161,0.00091850903,0.00058197015,0.0005264683,0.001238155,0.0018920207,0.0010154713,0.0053100237],"category_scores_gemma":[0.0059886724,0.00067756895,0.0009366103,0.00061352714,0.00068254356,0.0028561414,0.0018382128,0.0020621184,0.002110726],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075775123,0.00020825276,0.0047369096,0.00023143864,0.00025410153,0.0003046797,0.0003760059,0.55736125,0.024452282,0.020609412,0.01881814,0.37188992],"study_design_scores_gemma":[0.000011898952,0.000015744561,0.00009863886,0.0000036773665,0.000008432214,0.000022024635,0.000015867925,0.98871136,0.0041957656,0.005990189,0.00091961,0.0000067199144],"about_ca_topic_score_codex":0.009820026,"about_ca_topic_score_gemma":0.023237709,"teacher_disagreement_score":0.009820026,"about_ca_system_score_codex":0.00095164124,"about_ca_system_score_gemma":0.0019222924,"threshold_uncertainty_score":0.019525707},"labels":[],"label_agreement":null},{"id":"W4409160206","doi":"10.1007/978-3-031-88711-6_15","title":"Rank-Without-GPT: Building GPT-Independent Listwise Rerankers on Open-Source Large Language Models","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Regional Municipality of Waterloo; University of Waterloo","funders":"","keywords":"Computer science; Rank (graph theory); Open source; Artificial intelligence; Programming language; Natural language processing; Mathematics; Software","score_opus":0.01984332292328901,"score_gpt":0.2749590694641937,"score_spread":0.2551157465409047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409160206","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017284706,0.00187355,0.76825076,0.0006615541,0.00090302044,0.0005364853,0.011000973,0.1930232,0.00646572],"genre_scores_gemma":[0.092910126,0.0006499689,0.8330901,0.0006597906,0.0006152025,0.00069531077,0.044139378,0.008994575,0.018245587],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972535,0.0008283336,0.00018087155,0.00075971906,0.0006495951,0.00032803338],"domain_scores_gemma":[0.9934436,0.0032409343,0.00018009947,0.0018262551,0.0010670434,0.00024209508],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032081737,0.004201436,0.0027269998,0.0036818467,0.0014250652,0.002679325,0.0038375824,0.0033715728,0.024321882],"category_scores_gemma":[0.01600742,0.0017502942,0.0028598513,0.0029280249,0.00083805487,0.0058367862,0.0028504138,0.004752222,0.033015657],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005370517,0.0004188762,0.0012716524,0.00052285404,0.00039400003,0.0002680621,0.00015811752,0.051554706,0.006536636,0.0067645144,0.20094907,0.7306245],"study_design_scores_gemma":[0.00038962573,0.0002476125,0.0004992284,0.00005649378,0.00014177729,0.00017900534,0.00010362563,0.9404159,0.0072571216,0.032850746,0.017777465,0.00008129474],"about_ca_topic_score_codex":0.010463653,"about_ca_topic_score_gemma":0.039497875,"teacher_disagreement_score":0.024321882,"about_ca_system_score_codex":0.0011813994,"about_ca_system_score_gemma":0.0034541436,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4409166237","doi":"10.1007/978-3-031-88720-8_51","title":"Eval4RAG: Workshop on Evaluation of Retrieval-Augmented Generation Systems","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Artificial intelligence","score_opus":0.06232949037925545,"score_gpt":0.3027192913743151,"score_spread":0.24038980099505963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409166237","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.109155826,0.02924948,0.6485755,0.004566624,0.006053441,0.002956212,0.030464083,0.0866713,0.082307585],"genre_scores_gemma":[0.35229042,0.003339604,0.5032074,0.0016151316,0.00078814616,0.0015325821,0.06075885,0.009415038,0.067052744],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97908473,0.012582921,0.0007430185,0.0017335401,0.0050899563,0.00076577876],"domain_scores_gemma":[0.9875311,0.0061650095,0.00024223913,0.0028522417,0.0028590912,0.0003502263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016924154,0.0032258877,0.0026706643,0.0026073807,0.0009810925,0.004576709,0.0068859295,0.0035506375,0.039248716],"category_scores_gemma":[0.025803449,0.00092759525,0.0016112614,0.0019568661,0.0012227125,0.0050130617,0.0040445244,0.002346106,0.009421802],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006632835,0.0012225249,0.0019938643,0.002016799,0.00075501,0.00028233332,0.00037598537,0.041911785,0.015564366,0.01309644,0.20183307,0.71431506],"study_design_scores_gemma":[0.0026510553,0.0033787016,0.0054055257,0.0007002144,0.0008854275,0.00059706083,0.0005707527,0.6429775,0.07915612,0.037372023,0.22603028,0.0002754279],"about_ca_topic_score_codex":0.012069837,"about_ca_topic_score_gemma":0.010051667,"teacher_disagreement_score":0.039248716,"about_ca_system_score_codex":0.0026602333,"about_ca_system_score_gemma":0.002192268,"threshold_uncertainty_score":0.13129997},"labels":[],"label_agreement":null},{"id":"W4409167231","doi":"10.1007/978-3-031-88708-6_3","title":"Is Relevance Propagated from Retriever to Generator in RAG?","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Relevance (law); Generator (circuit theory); Labrador Retriever; Medicine; Physics; Political science; Law; Pathology","score_opus":0.01878913682042273,"score_gpt":0.24553504560915881,"score_spread":0.22674590878873607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409167231","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32735533,0.008237931,0.4775711,0.0164806,0.0031556939,0.0007404228,0.007971238,0.048225448,0.11026222],"genre_scores_gemma":[0.8929616,0.0016549054,0.05443299,0.0013902555,0.0016333868,0.000089570174,0.005498135,0.0054802108,0.03685905],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943896,0.0014962793,0.0002847231,0.001344436,0.0017952543,0.0006896785],"domain_scores_gemma":[0.9840919,0.0071633263,0.0006877075,0.0043753847,0.0030224305,0.0006592897],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050936067,0.0009815679,0.001472078,0.0031575179,0.0014345243,0.006322133,0.0024489346,0.0025292232,0.018174183],"category_scores_gemma":[0.0318363,0.0013203783,0.0009822232,0.0021341303,0.0016550259,0.015918907,0.0025640605,0.002532354,0.013588029],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028263677,0.0005548948,0.024970077,0.0017158089,0.00046816794,0.0022155894,0.0046739993,0.0090874685,0.1229016,0.12842242,0.122360095,0.57980347],"study_design_scores_gemma":[0.00051458646,0.0010237942,0.040731344,0.00040182806,0.0013582426,0.00412679,0.0028179768,0.15561453,0.18632953,0.4040999,0.20237248,0.0006090508],"about_ca_topic_score_codex":0.0045475317,"about_ca_topic_score_gemma":0.003103384,"teacher_disagreement_score":0.018174183,"about_ca_system_score_codex":0.0018886682,"about_ca_system_score_gemma":0.001821515,"threshold_uncertainty_score":0.060798645},"labels":[],"label_agreement":null},{"id":"W4409167270","doi":"10.1007/978-3-031-88708-6_16","title":"Lost but Not Only in the Middle","year":2025,"lang":"fr","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Universiteit van Amsterdam; Canadian Institute of Steel Construction","keywords":"Computer science","score_opus":0.05394133038696452,"score_gpt":0.2599273289173427,"score_spread":0.2059859985303782,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409167270","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027743888,0.040935047,0.021927534,0.0659425,0.031659342,0.00007333384,0.000767939,0.0012157781,0.83470416],"genre_scores_gemma":[0.011613828,0.006616006,0.0025100105,0.0065653715,0.003392333,0.000031673153,0.00024059784,0.0007996372,0.9682304],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993204,0.00011591211,0.00002820788,0.0001867542,0.0002446695,0.000104058636],"domain_scores_gemma":[0.9988207,0.00028887033,0.000046718153,0.0002214193,0.0003624507,0.00025987855],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009878573,0.0009121873,0.0011032646,0.0014169228,0.0037548072,0.013681819,0.0011418859,0.0022207047,0.13078322],"category_scores_gemma":[0.0037221648,0.0006564914,0.0008011925,0.0013887322,0.0033369975,0.022922518,0.004786941,0.00883525,0.09209767],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009955866,0.000040082345,0.00016792685,0.00026995444,0.000019713152,0.00017361104,0.0031896152,0.00013024523,0.0014402672,0.2769248,0.56729317,0.15025114],"study_design_scores_gemma":[0.0000026322166,0.00001015067,0.00006593842,0.00010194428,0.0000049415107,0.00014634323,0.0004169746,0.00004649818,0.00014848643,0.02274017,0.97630864,0.000007182426],"about_ca_topic_score_codex":0.00196084,"about_ca_topic_score_gemma":0.003238388,"teacher_disagreement_score":0.13078322,"about_ca_system_score_codex":0.0019478776,"about_ca_system_score_gemma":0.0012791789,"threshold_uncertainty_score":0.4375134},"labels":[],"label_agreement":null},{"id":"W4409168760","doi":"10.1145/3728373","title":"Recall, Robustness, and Lexicographic Evaluation","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Recommender Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Microsoft (Canada)","funders":"","keywords":"Lexicographical order; Robustness (evolution); Recall; Computer science; Artificial intelligence; Cognitive psychology; Mathematics; Psychology; Chemistry; Combinatorics","score_opus":0.05255388057789485,"score_gpt":0.2979719475362509,"score_spread":0.24541806695835605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409168760","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16495243,0.0054069064,0.7915114,0.0025635343,0.00028454335,0.0006626607,0.0014489199,0.00061614776,0.03255344],"genre_scores_gemma":[0.84028304,0.00089369074,0.15382238,0.00060245243,0.00041086,0.0006187239,0.0011767605,0.00018702912,0.0020050071],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9248484,0.043828435,0.006632861,0.0061926143,0.017272625,0.0012250507],"domain_scores_gemma":[0.6673087,0.25700557,0.027848737,0.03025647,0.015693244,0.0018872586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.053536803,0.0014979509,0.0019205065,0.009679002,0.0014949683,0.0059272046,0.001565361,0.0019623113,0.003206602],"category_scores_gemma":[0.27469733,0.00050745707,0.0016831437,0.008139462,0.0070576677,0.011884118,0.004184558,0.0022118187,0.0006979393],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020776577,0.0005038314,0.09188061,0.0020970646,0.0013566914,0.00036692148,0.0048069493,0.057431325,0.004201259,0.44007894,0.006346952,0.38885173],"study_design_scores_gemma":[0.00023139533,0.0020821628,0.043169305,0.0009665648,0.0005911299,0.0010390318,0.0021197035,0.13899225,0.010999877,0.78046095,0.01878065,0.0005670422],"about_ca_topic_score_codex":0.0016530567,"about_ca_topic_score_gemma":0.0013284549,"teacher_disagreement_score":0.053536803,"about_ca_system_score_codex":0.0030587048,"about_ca_system_score_gemma":0.001593243,"threshold_uncertainty_score":0.28313303},"labels":[],"label_agreement":null},{"id":"W4409183237","doi":"10.1007/978-3-031-88714-7_1","title":"exHarmony: Authorship and Citations for Benchmarking the Reviewer Assignment Problem","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Systems, Applications & Products in Data Processing (Canada)","funders":"","keywords":"Computer science; Benchmarking; Information retrieval; Management","score_opus":0.03858181707379498,"score_gpt":0.27900060615413147,"score_spread":0.24041878908033648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409183237","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27539364,0.013033265,0.56475157,0.0073900223,0.006143174,0.0009836036,0.047055874,0.028691767,0.056557126],"genre_scores_gemma":[0.59661084,0.0011795011,0.33721945,0.00045205475,0.0011944238,0.0007152364,0.041190643,0.0032247966,0.018213032],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9721121,0.014855402,0.0019469087,0.0032502522,0.006897643,0.00093775016],"domain_scores_gemma":[0.86578846,0.09027277,0.0050049596,0.02272382,0.012563571,0.0036463898],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.022588722,0.0016467167,0.0026513897,0.011177999,0.002591868,0.0056369537,0.004162424,0.00430251,0.0133250095],"category_scores_gemma":[0.14873025,0.0005817303,0.0011812904,0.012701409,0.0016634332,0.0074380725,0.004812278,0.00243384,0.0049097906],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001911836,0.0012508519,0.033925407,0.0018650334,0.00074955367,0.0002963611,0.0009635931,0.09808451,0.0023397845,0.073654935,0.24642085,0.5385374],"study_design_scores_gemma":[0.00056564395,0.00059638225,0.01143652,0.00036091328,0.0002608475,0.00050637114,0.0007210148,0.7256965,0.006878468,0.2027539,0.050055567,0.00016794987],"about_ca_topic_score_codex":0.0022065046,"about_ca_topic_score_gemma":0.0032672319,"teacher_disagreement_score":0.97741127,"about_ca_system_score_codex":0.0018718222,"about_ca_system_score_gemma":0.0032962114,"threshold_uncertainty_score":0.11946195},"labels":[],"label_agreement":null},{"id":"W4409183264","doi":"10.1007/978-3-031-88714-7_14","title":"The Impact of Incidental Multilingual Text on Cross-Lingual Transfer in Monolingual Retrieval","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Transfer (computing); Information retrieval; Speech recognition","score_opus":0.019360761337072593,"score_gpt":0.31301041212996555,"score_spread":0.29364965079289296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409183264","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9327972,0.0026126022,0.017475458,0.0006377086,0.00033885456,0.000082100305,0.00051058055,0.0011171036,0.044428386],"genre_scores_gemma":[0.9913244,0.00064785517,0.0028980495,0.00011784802,0.00018836389,0.00003402569,0.00048312268,0.00038868084,0.003917678],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964585,0.0018734712,0.00024982102,0.00052221434,0.00055998925,0.0003361112],"domain_scores_gemma":[0.941327,0.051533353,0.00080928113,0.002549009,0.0030042545,0.0007772081],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040170737,0.00055101427,0.0008354532,0.0011642838,0.0010812032,0.0036381623,0.0007431474,0.00096062326,0.010018106],"category_scores_gemma":[0.053015567,0.0004632313,0.00041811293,0.0014721889,0.0012404498,0.008215994,0.0034834193,0.0015515012,0.004549302],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007888349,0.0013516144,0.026034415,0.0013965961,0.00037403303,0.0011506282,0.0047694575,0.014102612,0.14398739,0.008049108,0.01053792,0.78035784],"study_design_scores_gemma":[0.00069955655,0.004531024,0.27980536,0.0005264149,0.0030152558,0.005267534,0.0140793,0.28405827,0.31768402,0.06774653,0.021958608,0.0006281891],"about_ca_topic_score_codex":0.0028386454,"about_ca_topic_score_gemma":0.0025161183,"teacher_disagreement_score":0.010018106,"about_ca_system_score_codex":0.00076551875,"about_ca_system_score_gemma":0.001000973,"threshold_uncertainty_score":0.033513904},"labels":[],"label_agreement":null},{"id":"W4409258158","doi":"10.1613/jair.1.16665","title":"Detecting AI-Generated Text: Factors Influencing Detectability with Current Methods","year":2025,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Ministère de la Défense Nationale","keywords":"Current (fluid); Computer science; Artificial intelligence; Physics","score_opus":0.21929592559695024,"score_gpt":0.49869199389354396,"score_spread":0.2793960682965937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409258158","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51312244,0.030633083,0.39471453,0.01596719,0.0017184159,0.0012763097,0.009720131,0.007247591,0.025600307],"genre_scores_gemma":[0.899339,0.003534572,0.08482188,0.00092937384,0.0010851311,0.0003628375,0.006872129,0.0005274373,0.0025275303],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9726078,0.014207584,0.0030549897,0.0043269694,0.0050525153,0.0007501737],"domain_scores_gemma":[0.63956136,0.30719298,0.0131320665,0.026097022,0.012191284,0.0018252778],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03233095,0.0014580908,0.0013896669,0.006893528,0.0021172261,0.008502508,0.0028195463,0.0029712033,0.0028554655],"category_scores_gemma":[0.23836783,0.00058817695,0.0011867638,0.0048800143,0.0022661013,0.014758926,0.0030547192,0.0038259393,0.0038357181],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017453433,0.0006932091,0.37044615,0.0022705938,0.00072095584,0.00047706702,0.004773808,0.019026505,0.007965226,0.012726883,0.040452875,0.5387014],"study_design_scores_gemma":[0.00022830436,0.0005967573,0.13821816,0.0012844198,0.0006804547,0.0025937285,0.0063814926,0.6894081,0.025108917,0.07747814,0.057681806,0.00033959327],"about_ca_topic_score_codex":0.0043298593,"about_ca_topic_score_gemma":0.004582524,"teacher_disagreement_score":0.03233095,"about_ca_system_score_codex":0.0014423748,"about_ca_system_score_gemma":0.0016111541,"threshold_uncertainty_score":0.17098445},"labels":[],"label_agreement":null},{"id":"W4409262394","doi":"10.1109/wacv61041.2025.00604","title":"Multi-Modal Large Language Model with RAG Strategies in Soccer Commentary Generation","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Modal; Artificial intelligence; Speech recognition; Linguistics; Natural language processing; Philosophy","score_opus":0.022807344741933976,"score_gpt":0.28917109403037455,"score_spread":0.2663637492884406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409262394","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06077787,0.002314648,0.924166,0.0015989296,0.00033151818,0.0002518728,0.000777789,0.005772394,0.0040089246],"genre_scores_gemma":[0.76319355,0.0006100876,0.22062373,0.0012786839,0.00043970882,0.00046689462,0.0024050425,0.0005866224,0.0103956815],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99890745,0.00057545886,0.00004222031,0.00025879775,0.00011166701,0.00010433495],"domain_scores_gemma":[0.99792033,0.0015499656,0.00007903041,0.00012871258,0.00023226581,0.00008974641],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018416563,0.0011092351,0.0009903867,0.00071703043,0.0005787712,0.000967549,0.0017626776,0.0014723599,0.0035932919],"category_scores_gemma":[0.0059914007,0.0004246101,0.0009722308,0.00052548765,0.0006114694,0.0016310216,0.0013934904,0.0020553155,0.0015579676],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001042183,0.000358881,0.002075736,0.00048701244,0.00022800772,0.00061277417,0.0008697382,0.49810997,0.015714135,0.019087229,0.021516265,0.43989813],"study_design_scores_gemma":[0.000027945249,0.00005267562,0.00011613191,0.000010680798,0.000018641062,0.00003085891,0.00005203542,0.9908077,0.0015549491,0.0061819647,0.0011320444,0.000014498496],"about_ca_topic_score_codex":0.0070455377,"about_ca_topic_score_gemma":0.010662567,"teacher_disagreement_score":0.0070455377,"about_ca_system_score_codex":0.0008079469,"about_ca_system_score_gemma":0.0011125554,"threshold_uncertainty_score":0.0140090585},"labels":[],"label_agreement":null},{"id":"W4409346487","doi":"10.1609/aaai.v39i2.32174","title":"Can Generative Models Improve Self-Supervised Representation Learning?","year":2025,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Vector Institute","funders":"","keywords":"Generative grammar; Computer science; Representation (politics); Generative model; Artificial intelligence; Machine learning","score_opus":0.07205294050913724,"score_gpt":0.30408855288021647,"score_spread":0.23203561237107923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409346487","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015295086,0.0012396867,0.97778136,0.0015109901,0.00014461864,0.000052259347,0.0001302445,0.0017857413,0.0020600709],"genre_scores_gemma":[0.63369054,0.0014587138,0.355221,0.0015332663,0.00043812377,0.0002443296,0.0013433963,0.00087028334,0.0052004056],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975721,0.0013086632,0.00008010473,0.0005134049,0.000379108,0.00014651069],"domain_scores_gemma":[0.9907782,0.005623824,0.00046676927,0.0020223372,0.0008502539,0.00025858945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00527504,0.0014985579,0.0014871145,0.0011567986,0.00061017286,0.0020866534,0.002512331,0.0023099259,0.0027995822],"category_scores_gemma":[0.019329581,0.0007651935,0.0013722138,0.0010579015,0.0019518592,0.004613387,0.0027132703,0.0036833948,0.0017468135],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020270715,0.0002958967,0.0032498348,0.0002954426,0.0002742928,0.00011710395,0.00039469887,0.48747203,0.0041138786,0.047618218,0.0139733935,0.4419925],"study_design_scores_gemma":[0.000014686559,0.0000397373,0.00013489238,0.000021072385,0.000011711424,0.00003104241,0.000021017911,0.9582214,0.0010504904,0.039033327,0.0014089104,0.000011811926],"about_ca_topic_score_codex":0.0032641788,"about_ca_topic_score_gemma":0.0041839783,"teacher_disagreement_score":0.00527504,"about_ca_system_score_codex":0.0011149368,"about_ca_system_score_gemma":0.0009824857,"threshold_uncertainty_score":0.027897418},"labels":[],"label_agreement":null},{"id":"W4409348661","doi":"10.1609/aaai.v39i27.35029","title":"J&amp;H: Evaluating the Robustness of Large Language Models Under Knowledge-Injection Attacks in Legal Domain","year":2025,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Robustness (evolution); Computer science; Computer security; Chemistry; Biochemistry","score_opus":0.12608340550399452,"score_gpt":0.3864990687183631,"score_spread":0.2604156632143686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409348661","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84738815,0.0014916083,0.12591523,0.0015360268,0.00043610128,0.0010965654,0.0034976385,0.013252314,0.005386269],"genre_scores_gemma":[0.9184914,0.00014175505,0.07607382,0.00043547046,0.000058928137,0.0002680238,0.0033833645,0.00031835993,0.0008288327],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97263116,0.013568765,0.0026070615,0.00381951,0.00624784,0.001125711],"domain_scores_gemma":[0.8271473,0.13525984,0.009048958,0.019833704,0.0056934557,0.0030167534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026228398,0.001694125,0.00095285196,0.003118367,0.0011463261,0.002817942,0.0028270264,0.0034778721,0.0018342119],"category_scores_gemma":[0.13080223,0.0007218606,0.0017127726,0.0012606364,0.003183741,0.007326817,0.0040066913,0.003698647,0.0007757114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008037885,0.003546554,0.115631245,0.0019338392,0.0026219506,0.0012328529,0.002995183,0.5490549,0.030523002,0.011997568,0.020387404,0.2520376],"study_design_scores_gemma":[0.00017294603,0.0009172425,0.008643008,0.000046803656,0.00016388873,0.00021792465,0.00034349618,0.9684803,0.013819602,0.0056074723,0.001496343,0.000091026195],"about_ca_topic_score_codex":0.015017259,"about_ca_topic_score_gemma":0.009502847,"teacher_disagreement_score":0.026228398,"about_ca_system_score_codex":0.002116651,"about_ca_system_score_gemma":0.0027506524,"threshold_uncertainty_score":0.13871068},"labels":[],"label_agreement":null},{"id":"W4409354062","doi":"10.2196/71721","title":"Authors’ Reply: Enhancing AI-Driven Medical Translations: Considerations for Language Concordance","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Education","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Concordance; Linguistics; Psychology; Natural language processing; Computer science; Medicine; Philosophy; Internal medicine","score_opus":0.017954926257992474,"score_gpt":0.37081960671953157,"score_spread":0.35286468046153907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409354062","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00028316674,0.00045596875,0.00049214344,0.9741538,0.023743246,0.000010681872,0.0001277254,0.000050234925,0.0006830015],"genre_scores_gemma":[0.00659739,0.00083335384,0.0018868946,0.9645908,0.021747567,0.00007650373,0.00013291345,0.000091956244,0.004042588],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9823293,0.006710854,0.0034876298,0.0015134937,0.0047739265,0.0011846896],"domain_scores_gemma":[0.8435478,0.092288226,0.0049463077,0.0038257984,0.051128976,0.004262886],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01895608,0.00062301586,0.0011464838,0.0015796113,0.0034271423,0.0049257022,0.0019257978,0.027140336,0.01977352],"category_scores_gemma":[0.20327125,0.0006894929,0.0013085428,0.0012224201,0.0033545676,0.0046260413,0.0031093026,0.026307032,0.0113777295],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007606965,0.000022663977,0.0005038652,0.00020923775,0.000038196384,0.0003175884,0.00074423745,0.00007585493,0.00025939007,0.0029972724,0.98343974,0.011315899],"study_design_scores_gemma":[0.00012672976,0.00004460375,0.0012133394,0.0008707996,0.000102755424,0.0013201303,0.0021583955,0.00064733933,0.0010684008,0.01200278,0.98030955,0.0001350817],"about_ca_topic_score_codex":0.00692642,"about_ca_topic_score_gemma":0.008801902,"teacher_disagreement_score":0.027140336,"about_ca_system_score_codex":0.003304919,"about_ca_system_score_gemma":0.008469327,"threshold_uncertainty_score":0.10025054},"labels":[],"label_agreement":null},{"id":"W4409361081","doi":"10.1609/aaai.v39i28.35203","title":"Stress-Testing of Multimodal Models in Medical Image-Based Report Generation","year":2025,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Fundação para a Ciência e a Tecnologia; Ministério da Ciência, Tecnologia e Ensino Superior; European Commission","keywords":"Stress testing (software); Computer science; Stress (linguistics); Artificial intelligence; Computer vision; Linguistics; Programming language; Philosophy","score_opus":0.11769881698210594,"score_gpt":0.33288271202371394,"score_spread":0.215183895041608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409361081","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.298064,0.00024592265,0.6920182,0.0019785059,0.00015985424,0.0002629325,0.00019202201,0.003512545,0.0035659582],"genre_scores_gemma":[0.93449557,0.00006489663,0.06419389,0.00020273613,0.00004413789,0.00012422154,0.00015523785,0.00022459749,0.00049474987],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9881052,0.0071461154,0.0005970603,0.0012141092,0.002481341,0.000456088],"domain_scores_gemma":[0.9062104,0.070538454,0.004325886,0.0121510215,0.0053812517,0.0013929701],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01568874,0.0008912881,0.0006323767,0.0007382428,0.00056416716,0.0028289938,0.002378462,0.0021354894,0.0035529344],"category_scores_gemma":[0.13948667,0.0006399833,0.0008124232,0.00044578093,0.0025313233,0.005305213,0.0033907099,0.0018327555,0.0009620388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044171666,0.0016080632,0.04590237,0.0005810865,0.0003534731,0.0015807776,0.0080623925,0.28517506,0.061946455,0.10486687,0.0062648035,0.47924152],"study_design_scores_gemma":[0.000058780268,0.0004305099,0.0019633938,0.000050074388,0.000042283766,0.00016919927,0.00034569376,0.9292371,0.024921725,0.04136897,0.0013529386,0.000059234335],"about_ca_topic_score_codex":0.0017175947,"about_ca_topic_score_gemma":0.0010335088,"teacher_disagreement_score":0.01568874,"about_ca_system_score_codex":0.0011521993,"about_ca_system_score_gemma":0.0013247448,"threshold_uncertainty_score":0.08297092},"labels":[],"label_agreement":null},{"id":"W4409364177","doi":"10.1609/aaai.v39i17.33989","title":"TimeCAP: Learning to Contextualize, Augment, and Predict Time Series Events with Large Language Model Agents","year":2025,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology; National Research Foundation of Korea; National Research Foundation","keywords":"Augment; Series (stratigraphy); Computer science; Linguistics; Philosophy; Biology","score_opus":0.034253247239254575,"score_gpt":0.2902851378484357,"score_spread":0.2560318906091811,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409364177","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033833124,0.0012498872,0.94599664,0.0008154284,0.00031790382,0.00016911254,0.0020067918,0.013234845,0.0023763825],"genre_scores_gemma":[0.45466065,0.0009882041,0.52970564,0.0009765124,0.00034896896,0.0005161205,0.0067546056,0.00064766116,0.005401657],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99948883,0.00016320928,0.000027387241,0.00020330094,0.00007955724,0.00003759086],"domain_scores_gemma":[0.99830437,0.0011122399,0.00014042249,0.00019001444,0.00017304625,0.00007991378],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011250676,0.0018132336,0.00066583714,0.00087389705,0.00047815003,0.0011166752,0.001413855,0.0009967164,0.002386753],"category_scores_gemma":[0.0049700513,0.0004906588,0.001288145,0.0008326974,0.0004561987,0.0027372325,0.0018123207,0.0025005224,0.0014205578],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005754271,0.0007434528,0.008652894,0.0005168556,0.00039482914,0.00062148646,0.00067160657,0.31049913,0.018947653,0.014729146,0.031441655,0.6122059],"study_design_scores_gemma":[0.000011137735,0.00005329177,0.00032308194,0.000018469946,0.000030924555,0.000035092977,0.000039342733,0.98767763,0.0022700445,0.0069118356,0.0026106471,0.000018542676],"about_ca_topic_score_codex":0.0072253854,"about_ca_topic_score_gemma":0.012785613,"teacher_disagreement_score":0.0072253854,"about_ca_system_score_codex":0.00058590976,"about_ca_system_score_gemma":0.0012372738,"threshold_uncertainty_score":0.014366686},"labels":[],"label_agreement":null},{"id":"W4409369201","doi":"10.1609/aaai.v39i1.32020","title":"ChemVLM: Exploring the Power of Multimodal Large Language Models in Chemistry Area","year":2025,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Power (physics); Computer science; Chemistry; Linguistics; Philosophy; Physics; Thermodynamics","score_opus":0.0846630502594492,"score_gpt":0.29566091103300757,"score_spread":0.21099786077355837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409369201","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22214508,0.013155772,0.5806561,0.0073690037,0.0014482094,0.0014062151,0.059124853,0.09246098,0.022233812],"genre_scores_gemma":[0.4777415,0.002712648,0.39781997,0.0028213975,0.00039812716,0.0015570936,0.10359769,0.003130578,0.010220998],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9986615,0.00061341195,0.00007665068,0.00039197886,0.00018218056,0.00007431577],"domain_scores_gemma":[0.9951762,0.0035905675,0.00014751448,0.00057122187,0.0003539512,0.0001605798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030709447,0.0024621508,0.00093493715,0.0019474605,0.00089513836,0.0021199016,0.002726026,0.0022532865,0.0063905967],"category_scores_gemma":[0.010738044,0.0006515051,0.0022790777,0.0012789207,0.000570807,0.0042168032,0.0025059031,0.003541608,0.0036757085],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009805865,0.0011751787,0.008672677,0.0017339556,0.00094162964,0.0006852472,0.00064526795,0.37839827,0.008280212,0.01605445,0.12629074,0.45614174],"study_design_scores_gemma":[0.00012959147,0.00014128152,0.00069186476,0.00006510373,0.00008189741,0.00010922597,0.00013246691,0.96983117,0.0033273464,0.012105628,0.0133360615,0.000048448168],"about_ca_topic_score_codex":0.023977246,"about_ca_topic_score_gemma":0.034413382,"teacher_disagreement_score":0.023977246,"about_ca_system_score_codex":0.002165279,"about_ca_system_score_gemma":0.0022907767,"threshold_uncertainty_score":0.04767537},"labels":[],"label_agreement":null},{"id":"W4409374936","doi":"10.21203/rs.3.rs-6411861/v1","title":"An Optimized Content Retriever from Web Articles using Large Language Models and FAISS Indexing","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Search engine indexing; Content (measure theory); Computer science; Information retrieval; Labrador Retriever; World Wide Web; Mathematics; Medicine","score_opus":0.19370522112219066,"score_gpt":0.4134479178797606,"score_spread":0.21974269675756994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409374936","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03973111,0.0019168709,0.8032958,0.0005491511,0.00045919168,0.0005345,0.007352306,0.14107627,0.005084854],"genre_scores_gemma":[0.16956353,0.0007778987,0.7854117,0.00030886242,0.00057666324,0.00044894242,0.019874549,0.004758713,0.018279],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99899656,0.00016167319,0.00011106434,0.0002699807,0.00032581826,0.00013481583],"domain_scores_gemma":[0.998156,0.0007557328,0.000087550434,0.00035502692,0.00053374644,0.00011186408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011665435,0.0016231982,0.0021812827,0.0049081924,0.0010408985,0.0026456523,0.0018327815,0.0013764127,0.012739046],"category_scores_gemma":[0.0041882577,0.0007238914,0.0016952254,0.003686203,0.0003649919,0.0032486073,0.0015207763,0.00094516465,0.016064148],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013817213,0.00036175264,0.0014379882,0.00041618638,0.0002543987,0.00027645312,0.00013762117,0.008568093,0.06803559,0.004036992,0.057597563,0.85749567],"study_design_scores_gemma":[0.0004219936,0.00044323146,0.0028138596,0.00003126223,0.0003517058,0.00065401837,0.00025984133,0.8674105,0.081798196,0.011035345,0.034641594,0.00013846405],"about_ca_topic_score_codex":0.007834755,"about_ca_topic_score_gemma":0.013184067,"teacher_disagreement_score":0.012739046,"about_ca_system_score_codex":0.00096865266,"about_ca_system_score_gemma":0.0019775925,"threshold_uncertainty_score":0.042616308},"labels":[],"label_agreement":null},{"id":"W4409422935","doi":"10.1016/j.cjca.2025.04.009","title":"Editorial Commentary to Beyond Assistance: Are Large Language Models Ready for Autonomous Electrocardiogram Interpretation?","year":2025,"lang":"en","type":"editorial","venue":"Canadian Journal of Cardiology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Medicine; Interpretation (philosophy); Natural language processing; Linguistics","score_opus":0.00770929327199492,"score_gpt":0.2651341505166486,"score_spread":0.2574248572446537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409422935","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000045490186,0.0033902926,0.00024062517,0.18445794,0.8106906,0.00002345588,0.00016989825,0.000070450325,0.0009112526],"genre_scores_gemma":[0.00065758824,0.0021001636,0.0002059092,0.0748247,0.9182469,0.000046105728,0.00006038405,0.00005792083,0.003800333],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9921996,0.0018365006,0.001034073,0.001082789,0.0033149116,0.0005321274],"domain_scores_gemma":[0.93130887,0.037699725,0.0029202958,0.001454248,0.020997921,0.005619015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011432342,0.0028632581,0.0040879142,0.0034685258,0.004078558,0.008700187,0.0033928873,0.029268987,0.015255919],"category_scores_gemma":[0.07160516,0.0013575585,0.0026439417,0.0015361294,0.0034442984,0.003592295,0.0017872875,0.026962498,0.010159789],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031015505,0.0000046980417,0.000021462249,0.000091731294,0.000015217997,0.000055011253,0.000011093397,0.000013334019,0.000015917949,0.0001921708,0.9979395,0.0016088915],"study_design_scores_gemma":[0.00026804628,0.00003624461,0.00037960763,0.0010970022,0.00016226816,0.0003438918,0.000109311775,0.00065708143,0.00013174379,0.0033447186,0.99341476,0.0000553804],"about_ca_topic_score_codex":0.0076472447,"about_ca_topic_score_gemma":0.013471457,"teacher_disagreement_score":0.029268987,"about_ca_system_score_codex":0.0042044735,"about_ca_system_score_gemma":0.0062567666,"threshold_uncertainty_score":0.060460687},"labels":[],"label_agreement":null},{"id":"W4409526936","doi":"10.5430/wjel.v15n5p272","title":"Optimizing Automated Essay Scoring: A Comparative Study of Machine Learning Approaches with a Focus on Ensemble Methods","year":2025,"lang":"en","type":"article","venue":"World Journal of English Language","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Focus (optics); Computer science; Ensemble learning; Artificial intelligence; Machine learning","score_opus":0.04555861244710417,"score_gpt":0.3266675701460139,"score_spread":0.2811089576989097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409526936","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31864375,0.008582069,0.66199,0.0012725511,0.00021980173,0.0002500116,0.00016589131,0.0016943853,0.007181629],"genre_scores_gemma":[0.86060715,0.0014697395,0.1353587,0.00018174798,0.00014118645,0.00014746653,0.0003143471,0.00014746908,0.0016322949],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99350643,0.0039241305,0.00042576715,0.000713875,0.0012058727,0.00022389137],"domain_scores_gemma":[0.97545016,0.01845193,0.0009604328,0.0015558812,0.0032438408,0.0003377934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010692603,0.0013142982,0.001435811,0.0016001171,0.00039592414,0.0015365989,0.0012401426,0.0011431446,0.0006233429],"category_scores_gemma":[0.027074855,0.00041588984,0.00070421933,0.0014634569,0.00043933193,0.001818562,0.0013509651,0.0015567555,0.0002811953],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031974842,0.00022386035,0.01174231,0.00022947151,0.00034812288,0.000043778535,0.0002677425,0.48954135,0.00148484,0.0024721418,0.0014134189,0.49191317],"study_design_scores_gemma":[0.000014264712,0.00026209425,0.003057203,0.000035372905,0.000044407105,0.000026743583,0.00007879845,0.9924591,0.0012934855,0.0016837483,0.0010266915,0.000018113189],"about_ca_topic_score_codex":0.0032275028,"about_ca_topic_score_gemma":0.003170123,"teacher_disagreement_score":0.010692603,"about_ca_system_score_codex":0.00084249704,"about_ca_system_score_gemma":0.001087478,"threshold_uncertainty_score":0.056548536},"labels":[],"label_agreement":null},{"id":"W4409561214","doi":"10.1109/ieeedata.2025.3562173","title":"Descriptor: Open-Domain Long-Form Context-Aware Question-Answering Dataset (DragonVerseQA)","year":2025,"lang":"en","type":"article","venue":"IEEE data descriptions.","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Open domain; Question answering; Computer science; Context (archaeology); Information retrieval; Domain (mathematical analysis); Closed-ended question; Artificial intelligence; Mathematics; Geography; Statistics","score_opus":0.08150532100520148,"score_gpt":0.3335827149188032,"score_spread":0.25207739391360173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409561214","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007024798,0.00096754666,0.0026418508,0.0007568296,0.00033050842,0.00046953937,0.97436666,0.0071002403,0.0063421256],"genre_scores_gemma":[0.0064051202,0.00010504085,0.004774791,0.00023573323,0.000042408737,0.00039852835,0.985827,0.00017876946,0.0020325878],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982889,0.00040215426,0.00019304904,0.0004922407,0.0004517068,0.0001718596],"domain_scores_gemma":[0.9969733,0.00087709667,0.00022443064,0.0006839153,0.00085652305,0.00038481614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015340347,0.0022249233,0.00096359605,0.0027405522,0.0013962834,0.0018611902,0.0022522053,0.003498199,0.025530184],"category_scores_gemma":[0.008327253,0.0003799568,0.0011171896,0.0028044854,0.00069751305,0.0024749388,0.0032649082,0.002124398,0.031704534],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027398247,0.0001552152,0.0015700543,0.0015057107,0.00004675791,0.00017657242,0.00041454122,0.0010968794,0.0022796397,0.0013146393,0.97111017,0.02005582],"study_design_scores_gemma":[0.00032561278,0.00017851197,0.008798587,0.00039252886,0.000048916292,0.000303409,0.0011542931,0.008429121,0.0037351795,0.0034252554,0.9731026,0.00010589615],"about_ca_topic_score_codex":0.022111062,"about_ca_topic_score_gemma":0.04576639,"teacher_disagreement_score":0.025530184,"about_ca_system_score_codex":0.0013903109,"about_ca_system_score_gemma":0.0018003887,"threshold_uncertainty_score":0.08540702},"labels":[],"label_agreement":null},{"id":"W4409572435","doi":"10.1016/j.eswa.2025.127612","title":"Fact retrieval from knowledge graphs through semantic and contextual attention","year":2025,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Saint Mary's University; Dalhousie University; Cape Breton University","funders":"Faculty of Graduate Studies and Research, University of Alberta; Natural Sciences and Engineering Research Council of Canada; Southern Methodist University; Saint Mary’s University","keywords":"Computer science; Knowledge graph; Semantic memory; Information retrieval; Natural language processing; Artificial intelligence; Cognition; Psychology","score_opus":0.02090410090579666,"score_gpt":0.2828838789322674,"score_spread":0.26197977802647077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409572435","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07371933,0.0052207736,0.8909152,0.0008650264,0.00021917104,0.00034422305,0.004498389,0.01765089,0.00656695],"genre_scores_gemma":[0.4511371,0.0024141085,0.5218453,0.0005754375,0.00025556132,0.00020015247,0.018329876,0.00074736576,0.0044950168],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989182,0.00022087566,0.000072202725,0.00040139852,0.00029656437,0.00009092684],"domain_scores_gemma":[0.99781775,0.00111337,0.00016662035,0.00048508923,0.0003250842,0.00009200008],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009439155,0.00147254,0.001047641,0.006742939,0.0008340465,0.001915725,0.001511363,0.0010095835,0.0028591035],"category_scores_gemma":[0.00624625,0.0004734283,0.0011801893,0.0047319746,0.00071872293,0.0054623745,0.002624617,0.0011592535,0.0014299245],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051632733,0.00030669116,0.0041825133,0.0010860538,0.00028638754,0.0005895783,0.0010478998,0.04067267,0.033392556,0.017741911,0.046303574,0.8538738],"study_design_scores_gemma":[0.00016530657,0.00041824413,0.009109776,0.00018134435,0.00058664393,0.0010330098,0.0011446428,0.74891865,0.04039369,0.1307901,0.06708404,0.00017459819],"about_ca_topic_score_codex":0.0144943,"about_ca_topic_score_gemma":0.029086443,"teacher_disagreement_score":0.0144943,"about_ca_system_score_codex":0.001119919,"about_ca_system_score_gemma":0.0014496689,"threshold_uncertainty_score":0.028819859},"labels":[],"label_agreement":null},{"id":"W4409605519","doi":"10.2196/71687","title":"Detecting Redundant Health Survey Questions by Using Language-Agnostic Bidirectional Encoder Representations From Transformers Sentence Embedding: Algorithm Development Study","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Computer science; Embedding; Sentence; Natural language processing; Artificial intelligence; Data science; World Wide Web","score_opus":0.034425854729574636,"score_gpt":0.37165474721816316,"score_spread":0.3372288924885885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409605519","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.091631815,0.00090669916,0.8971043,0.00068877556,0.00019148792,0.0006013443,0.0008776479,0.0062905294,0.0017073806],"genre_scores_gemma":[0.39936206,0.00047142798,0.58703333,0.0005008743,0.00015790708,0.0008820177,0.0068568233,0.0002985753,0.004436999],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985915,0.00046291543,0.00013980774,0.00043244674,0.00023220111,0.00014118433],"domain_scores_gemma":[0.9963709,0.0022671842,0.00022473269,0.00029167367,0.0007299423,0.00011546755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024940348,0.001445954,0.0010642421,0.0013195907,0.0005209426,0.0012252185,0.0017606826,0.0014449932,0.003900858],"category_scores_gemma":[0.008799331,0.00043044856,0.0009825446,0.0009795658,0.00053142256,0.0021836513,0.0016088195,0.0018491563,0.0019476227],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049639214,0.00040433,0.0065013114,0.0003275106,0.00016745577,0.0003349713,0.00035805177,0.054956473,0.010971194,0.003943852,0.010427361,0.9111111],"study_design_scores_gemma":[0.000049355385,0.00011966517,0.00083976204,0.000026335629,0.00005548974,0.00016036473,0.00015870167,0.9895704,0.0044357223,0.00282416,0.0017461947,0.000013922043],"about_ca_topic_score_codex":0.00548178,"about_ca_topic_score_gemma":0.0065852692,"teacher_disagreement_score":0.00548178,"about_ca_system_score_codex":0.00096794555,"about_ca_system_score_gemma":0.002219072,"threshold_uncertainty_score":0.013189912},"labels":[],"label_agreement":null},{"id":"W4409637115","doi":"10.1007/978-3-031-90167-6_16","title":"Future Sight: Fine-Tuning Language Models for Dynamic Story Generation","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Sight; Programming language; Astronomy; Physics","score_opus":0.018607415969301735,"score_gpt":0.2512150235771715,"score_spread":0.23260760760786978,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409637115","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012926329,0.00052010105,0.9177709,0.00039207598,0.0003778632,0.00012523484,0.0013629085,0.061424393,0.0051002316],"genre_scores_gemma":[0.27580526,0.00040838163,0.69427896,0.00054089905,0.00020273292,0.00038045467,0.0061240685,0.007149581,0.015109675],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99928004,0.00019856135,0.00004400491,0.00026855693,0.00014218634,0.00006670759],"domain_scores_gemma":[0.9982262,0.0010185343,0.000053746102,0.00034293806,0.00024790235,0.00011060309],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012152513,0.0012688537,0.00088711036,0.0006178375,0.00046642774,0.0020300644,0.0030361156,0.0016991291,0.022982918],"category_scores_gemma":[0.0060158265,0.0008544661,0.0010860059,0.00047452038,0.00049473287,0.0047228676,0.0022253203,0.0024293037,0.009167928],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010374135,0.00038355135,0.0011382138,0.00042275636,0.00016879005,0.00020327778,0.00045694868,0.10449736,0.029244246,0.021437846,0.07782082,0.7631888],"study_design_scores_gemma":[0.000073540315,0.000048498267,0.0001106147,0.00001506968,0.000030358067,0.000056265846,0.00006579838,0.9677559,0.0073930095,0.015323983,0.009107485,0.000019417048],"about_ca_topic_score_codex":0.004048639,"about_ca_topic_score_gemma":0.0070475964,"teacher_disagreement_score":0.022982918,"about_ca_system_score_codex":0.0008411955,"about_ca_system_score_gemma":0.0008112951,"threshold_uncertainty_score":0.07688552},"labels":[],"label_agreement":null},{"id":"W4409638536","doi":"10.1016/j.ijmedinf.2025.105942","title":"Ontology accelerates few-shot learning capability of large language model: A study in extraction of drug efficacy in a rare pediatric epilepsy","year":2025,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"National Center for Advancing Translational Sciences; National Institute of Biomedical Imaging and Bioengineering; Clinical and Translational Science Collaborative of Cleveland, School of Medicine, Case Western Reserve University; National Institutes of Health; Eisai; Dravet Syndrome Foundation; National Institute on Aging; Patient-Centered Outcomes Research Institute; National Institute on Drug Abuse; Epilepsy Foundation; U.S. Department of Defense","keywords":"Ontology; Epilepsy; Computer science; Antiepileptic drug; Drug; One shot; Natural language processing; Shot (pellet); Extraction (chemistry); Artificial intelligence; Information extraction; Machine learning; Medicine; Pharmacology","score_opus":0.024321061583311845,"score_gpt":0.3675098786225851,"score_spread":0.3431888170392733,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409638536","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.72190607,0.011917579,0.24350159,0.002169182,0.00046475948,0.0005092115,0.008696636,0.0064490433,0.0043858527],"genre_scores_gemma":[0.8191145,0.0019262127,0.15770614,0.0006363371,0.00018113636,0.00024058217,0.018172717,0.00019080701,0.0018315822],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988738,0.00034299525,0.00013728712,0.00036226158,0.00019682602,0.000086839354],"domain_scores_gemma":[0.9959727,0.0031259311,0.00020221609,0.00023249061,0.00034136418,0.00012535592],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019642864,0.0010021897,0.00071410456,0.0024607691,0.00052854035,0.00074399175,0.00089270144,0.001034278,0.0010470508],"category_scores_gemma":[0.006491667,0.00020754225,0.001453132,0.0015836932,0.00039696932,0.0017232619,0.0008144257,0.0010548864,0.0003435552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001665893,0.0015387697,0.038193475,0.0016618615,0.0007909343,0.0016445875,0.00047585808,0.13512811,0.017311845,0.0036296803,0.020319577,0.7776395],"study_design_scores_gemma":[0.0001299535,0.0004429746,0.008534415,0.00006798083,0.00031986754,0.00039635456,0.0002683796,0.96769834,0.0104029365,0.004258833,0.0074290517,0.000050984574],"about_ca_topic_score_codex":0.009276121,"about_ca_topic_score_gemma":0.012055347,"teacher_disagreement_score":0.009276121,"about_ca_system_score_codex":0.0009551998,"about_ca_system_score_gemma":0.001691995,"threshold_uncertainty_score":0.01844424},"labels":[],"label_agreement":null},{"id":"W4409710058","doi":"10.1101/2025.04.22.25326190","title":"ALPaCA: Adapting Llama for Pathology Context Analysis to enable slide-level question answering","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Context (archaeology); Question answering; Computer science; Pathology; Medicine; Information retrieval; Geography; Archaeology","score_opus":0.06290461622411195,"score_gpt":0.31104258771240734,"score_spread":0.2481379714882954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409710058","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04920864,0.0036092289,0.67770356,0.001177471,0.00069331203,0.0009928623,0.00661805,0.25423723,0.005759655],"genre_scores_gemma":[0.33865312,0.00072977715,0.62164074,0.003084284,0.00026508904,0.0017871485,0.021093791,0.004096926,0.008649099],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99806184,0.0005761936,0.00011176129,0.0008305081,0.0002717894,0.00014793643],"domain_scores_gemma":[0.9966497,0.0021036884,0.00014978713,0.00039953706,0.00049758534,0.00019974826],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024399038,0.0024189062,0.00094189146,0.0015961499,0.0005114156,0.0019248959,0.003352416,0.003012011,0.017493552],"category_scores_gemma":[0.011397264,0.000751941,0.0018784656,0.00057788077,0.0008321148,0.0029265606,0.00377061,0.0024354104,0.008427621],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012128681,0.0006650633,0.008540631,0.0023997116,0.00041821547,0.0009487948,0.0014073845,0.05772006,0.05630648,0.005614028,0.10289115,0.7618756],"study_design_scores_gemma":[0.00016074834,0.00042584792,0.0019799916,0.00019782175,0.000095576266,0.0004599484,0.0005051924,0.9159752,0.029349068,0.010252644,0.04049392,0.00010406485],"about_ca_topic_score_codex":0.0057772794,"about_ca_topic_score_gemma":0.010043438,"teacher_disagreement_score":0.017493552,"about_ca_system_score_codex":0.0010997846,"about_ca_system_score_gemma":0.0016079145,"threshold_uncertainty_score":0.058521748},"labels":[],"label_agreement":null},{"id":"W4409838812","doi":"10.1007/s10845-025-02614-4","title":"Gen-JEMA: enhanced explainability using generative joint embedding multimodal alignment for monitoring directed energy deposition","year":2025,"lang":"en","type":"article","venue":"Journal of Intelligent Manufacturing","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"European Regional Development Fund; Fundação para a Ciência e a Tecnologia","keywords":"Embedding; Joint (building); Deposition (geology); Computer science; Energy (signal processing); Materials science; Artificial intelligence; Engineering; Geology; Structural engineering; Physics","score_opus":0.03275531907417412,"score_gpt":0.30405742706484384,"score_spread":0.2713021079906697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409838812","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03190248,0.00021361114,0.9647669,0.00015580875,0.000033463642,0.000022901237,0.00013710924,0.0017299154,0.0010377602],"genre_scores_gemma":[0.75115556,0.00019261826,0.2435981,0.00023915921,0.000039927094,0.00011859354,0.0008293336,0.00057770795,0.0032489332],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99967766,0.00009273232,0.000011776894,0.000112377034,0.00007508912,0.000030327146],"domain_scores_gemma":[0.9993606,0.0003489598,0.000059305876,0.00011217588,0.00009178302,0.000027213882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007616527,0.00097018905,0.00052166067,0.00044728472,0.0002364688,0.0007318918,0.0010011466,0.0007653433,0.0020469518],"category_scores_gemma":[0.002356515,0.00042426639,0.00092656404,0.00034686996,0.00054499216,0.0011754978,0.0014650187,0.0013646835,0.00047976073],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015798131,0.00010337267,0.0017615729,0.00009228143,0.00009352873,0.00015367298,0.00014242604,0.8222536,0.023207255,0.00855747,0.0015699838,0.14190698],"study_design_scores_gemma":[0.000002998393,0.00001563417,0.0001410048,0.0000028232944,0.000004976669,0.000014044952,0.0000054490115,0.9938439,0.0029268805,0.0026934657,0.00034337377,0.000005426422],"about_ca_topic_score_codex":0.0027177408,"about_ca_topic_score_gemma":0.0039753234,"teacher_disagreement_score":0.0027177408,"about_ca_system_score_codex":0.00046881288,"about_ca_system_score_gemma":0.00060075417,"threshold_uncertainty_score":0.0068476796},"labels":[],"label_agreement":null},{"id":"W4409857686","doi":"10.1002/mef2.70019","title":"The Use of Large Language Models and Their Association With Enhanced Impact in Biomedical Research and Beyond","year":2025,"lang":"en","type":"article","venue":"MedComm – Future Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Impact","funders":"National Key Research and Development Program of China; National Natural Science Foundation of China","keywords":"Association (psychology); Computer science; Psychology; Data science; Psychotherapist","score_opus":0.029501484144652833,"score_gpt":0.3417388786316627,"score_spread":0.3122373944870099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409857686","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.893062,0.0055302754,0.07376672,0.0067872857,0.00035071414,0.00015678207,0.00411772,0.0021199984,0.014108606],"genre_scores_gemma":[0.98490924,0.00066972536,0.011489815,0.00018893262,0.00022987628,0.00007285919,0.0011884973,0.00026272182,0.0009883674],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98206306,0.011629172,0.0015561136,0.0017063677,0.0025209363,0.0005243386],"domain_scores_gemma":[0.6252857,0.3025851,0.03932613,0.016632093,0.013168155,0.0030028636],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.023257898,0.0006319366,0.0006146052,0.008567378,0.0010757836,0.007066973,0.0007075894,0.0009552663,0.0037472092],"category_scores_gemma":[0.18795778,0.0003463383,0.0007705865,0.008548617,0.001600181,0.0059550647,0.0031648956,0.0019597402,0.0013833592],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012391809,0.00041773383,0.6663148,0.0019003999,0.00080326037,0.0011074236,0.012587296,0.012186854,0.009603015,0.018772371,0.012105157,0.26296255],"study_design_scores_gemma":[0.00012460313,0.00063515874,0.59442484,0.0016481003,0.0011196783,0.0028634856,0.010547464,0.21264459,0.015892664,0.11370686,0.04590035,0.0004922376],"about_ca_topic_score_codex":0.001638743,"about_ca_topic_score_gemma":0.0021529212,"teacher_disagreement_score":0.9914326,"about_ca_system_score_codex":0.0012274062,"about_ca_system_score_gemma":0.001742856,"threshold_uncertainty_score":0.12300098},"labels":[],"label_agreement":null},{"id":"W4409963278","doi":"10.1098/rsos.241776","title":"Generalization bias in large language model summarization of scientific research","year":2025,"lang":"en","type":"article","venue":"Royal Society Open Science","topic":"Topic Modeling","field":"Computer Science","cited_by":84,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Automatic summarization; Generalization; Computer science; Natural language processing; Artificial intelligence; Language model; Epistemology; Philosophy","score_opus":0.09027618539005793,"score_gpt":0.3897188774648272,"score_spread":0.29944269207476926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409963278","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48058197,0.0070134695,0.47241962,0.010942033,0.00082261674,0.0011774448,0.005333244,0.01394265,0.0077669797],"genre_scores_gemma":[0.88897246,0.0007092417,0.101043575,0.0019157145,0.0002840187,0.00065705297,0.004500854,0.00065827195,0.0012587146],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9282089,0.059129275,0.0037770462,0.004743576,0.0036859885,0.00045515335],"domain_scores_gemma":[0.5770692,0.37966502,0.018250307,0.016199464,0.0077161854,0.0010998662],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.079010956,0.0012051126,0.0012678287,0.0028859028,0.00083801855,0.0038754921,0.001884103,0.0016548133,0.0019648662],"category_scores_gemma":[0.30427298,0.00061904103,0.0013541534,0.0019186895,0.0012125307,0.0059645707,0.003116737,0.0026538216,0.0008025205],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003918729,0.0005375172,0.12155785,0.007417943,0.0033136026,0.0012001302,0.0229698,0.17392144,0.01541438,0.027575202,0.042903427,0.57926995],"study_design_scores_gemma":[0.00053478486,0.0010726112,0.030217495,0.0014031924,0.0009868818,0.0005581825,0.0037593732,0.8078074,0.015352414,0.10780269,0.030061573,0.00044341822],"about_ca_topic_score_codex":0.0044971756,"about_ca_topic_score_gemma":0.007212297,"teacher_disagreement_score":0.92098904,"about_ca_system_score_codex":0.0018902473,"about_ca_system_score_gemma":0.0023102928,"threshold_uncertainty_score":0.41785485},"labels":[],"label_agreement":null},{"id":"W4409972601","doi":"10.1002/ail2.122","title":"A Few‐Shot Learning Approach for a Multilingual Agro‐Information Question Answering System","year":2025,"lang":"en","type":"article","venue":"Applied AI Letters","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Australian Bird Study Association; Foreign, Commonwealth and Development Office; International Development Research Centre","keywords":"Question answering; Computer science; One shot; Information retrieval; Shot (pellet); Natural language processing; Artificial intelligence; Engineering","score_opus":0.01099179054557923,"score_gpt":0.24223061526669915,"score_spread":0.23123882472111992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409972601","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19200537,0.0005759082,0.77283937,0.0014181184,0.0001862119,0.0007162036,0.0016691408,0.027286153,0.0033035174],"genre_scores_gemma":[0.555169,0.00009125209,0.43576074,0.00066710927,0.000069225105,0.00032593167,0.0036784085,0.00024416833,0.003994158],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998089,0.0006539696,0.00011667304,0.00077139534,0.00027128018,0.00009764886],"domain_scores_gemma":[0.9959413,0.002513589,0.00011230798,0.00038298336,0.0008193353,0.00023052937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032321343,0.0007332273,0.0007959339,0.0011673649,0.00097379007,0.0010317144,0.0022229098,0.0019928776,0.0043062684],"category_scores_gemma":[0.0088329865,0.00037362284,0.000777835,0.00060403947,0.000728447,0.0029670098,0.0019928424,0.0021661464,0.0016726301],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012777217,0.0015235344,0.0077572465,0.0010099434,0.00017031391,0.0010560394,0.0033092166,0.06500414,0.10958341,0.0068591624,0.016207797,0.7862416],"study_design_scores_gemma":[0.000047173024,0.0003037086,0.0017705413,0.000023899764,0.00003826306,0.00020816454,0.0005297704,0.95621145,0.02734663,0.0058923364,0.007586458,0.000041528532],"about_ca_topic_score_codex":0.007565625,"about_ca_topic_score_gemma":0.010568913,"teacher_disagreement_score":0.007565625,"about_ca_system_score_codex":0.0014850682,"about_ca_system_score_gemma":0.0011991457,"threshold_uncertainty_score":0.01709336},"labels":[],"label_agreement":null},{"id":"W4410004709","doi":"10.21449/ijate.1602294","title":"A review of automatic item generation techniques leveraging large language models","year":2025,"lang":"en","type":"review","venue":"International Journal of Assessment Tools in Education","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Data science","score_opus":0.07488583176221998,"score_gpt":0.4522201303716182,"score_spread":0.3773342986093982,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410004709","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013545714,0.9732526,0.02172538,0.00077938533,0.00021987618,0.00027109016,0.0003009002,0.00032981997,0.0017663927],"genre_scores_gemma":[0.02001863,0.90771514,0.068402626,0.0007257974,0.0004520812,0.0008182878,0.0009832553,0.0001278529,0.00075633824],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9903853,0.005451645,0.0013158119,0.0008204722,0.001901669,0.0001251643],"domain_scores_gemma":[0.9077246,0.083499305,0.002517563,0.0015544137,0.0044662147,0.0002378525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017555421,0.0022525801,0.0030628778,0.010652651,0.00052805996,0.002730724,0.0026124674,0.0012970417,0.004696497],"category_scores_gemma":[0.06382826,0.001179525,0.0033540525,0.008911084,0.00083107356,0.004513082,0.001342273,0.001481896,0.0026667858],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007061646,0.000084986175,0.0009808944,0.049854238,0.0005192439,0.000057518944,0.00030022569,0.001150876,0.00033162968,0.0015638396,0.0057237176,0.9393622],"study_design_scores_gemma":[0.00031850257,0.0013983275,0.026676567,0.18947968,0.00941641,0.0027341084,0.0017268696,0.019248417,0.005946143,0.026285151,0.71619064,0.00057909195],"about_ca_topic_score_codex":0.003400883,"about_ca_topic_score_gemma":0.0049482863,"teacher_disagreement_score":0.017555421,"about_ca_system_score_codex":0.001147437,"about_ca_system_score_gemma":0.004008513,"threshold_uncertainty_score":0.092842996},"labels":[],"label_agreement":null},{"id":"W4410041388","doi":"10.1007/978-3-031-86623-4_17","title":"Harnessing Pre-trained Language Models for Efficient Move Recognition in Biomedical Abstracts","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Information retrieval","score_opus":0.04928139306695458,"score_gpt":0.3130020145758459,"score_spread":0.26372062150889136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410041388","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021985635,0.0027878713,0.9350655,0.0009508712,0.0008421983,0.00018606317,0.0042974856,0.028237121,0.005647346],"genre_scores_gemma":[0.32212362,0.0042447145,0.62004924,0.0010263069,0.0014090776,0.00060400454,0.020294627,0.002242164,0.028006276],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99924636,0.00017041282,0.000063726366,0.00024385913,0.00016452221,0.000111128255],"domain_scores_gemma":[0.9984205,0.0008237891,0.000091541675,0.00018377547,0.00040343383,0.000076930715],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001125224,0.0018076985,0.0010683541,0.0020551141,0.00046932552,0.0018777533,0.0014326676,0.0014811701,0.0064092996],"category_scores_gemma":[0.0036758396,0.00054818,0.0017054563,0.0019663163,0.0003117744,0.0022209785,0.0012908077,0.001957695,0.013796254],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036824433,0.00016153742,0.0013213535,0.00034482862,0.0001528191,0.00021739268,0.00015881204,0.030447265,0.043914802,0.0036203547,0.040457275,0.8788353],"study_design_scores_gemma":[0.000026374815,0.00013008314,0.0012542388,0.00004796777,0.0001286888,0.0002467753,0.00014252735,0.9414646,0.033668894,0.008239058,0.0145926,0.000058173886],"about_ca_topic_score_codex":0.010130603,"about_ca_topic_score_gemma":0.013146167,"teacher_disagreement_score":0.010130603,"about_ca_system_score_codex":0.00073642255,"about_ca_system_score_gemma":0.0014863542,"threshold_uncertainty_score":0.021441221},"labels":[],"label_agreement":null},{"id":"W4410050200","doi":"10.1016/j.jallcom.2025.180709","title":"Data-driven explainable machine learning approaches for predicting hydrogen adsorption in porous crystalline materials","year":2025,"lang":"en","type":"article","venue":"Journal of Alloys and Compounds","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Adsorption; Porosity; Materials science; Porous medium; Hydrogen; Computer science; Chemistry; Physical chemistry; Composite material; Organic chemistry","score_opus":0.048832124473860754,"score_gpt":0.2627534371772068,"score_spread":0.21392131270334602,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410050200","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42159212,0.0016415572,0.5719744,0.0011908576,0.00008411841,0.00009098404,0.0013543098,0.0007629055,0.0013087355],"genre_scores_gemma":[0.9723126,0.00048707228,0.025146384,0.00004907686,0.000091562084,0.0001165636,0.00076354336,0.00005916594,0.00097402575],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99978226,0.00009363214,0.00001619191,0.000048648875,0.000034457466,0.000024797062],"domain_scores_gemma":[0.99604225,0.0033774574,0.00020318774,0.00013556225,0.00018603032,0.000055542183],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014100099,0.0005150199,0.0006482997,0.0012226233,0.00032401207,0.0007354821,0.0012064127,0.001063792,0.00092446705],"category_scores_gemma":[0.0050196443,0.000493914,0.0009327769,0.00079967047,0.00051584776,0.0014651003,0.0005875766,0.0009463179,0.00014207877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000057758425,0.000043993798,0.0020386067,0.000059121532,0.00006926756,0.0000347316,0.000040610623,0.98128307,0.0008136157,0.0065330854,0.00037551436,0.00865058],"study_design_scores_gemma":[0.0000020003145,0.000002563052,0.000099685065,8.330476e-7,0.0000031385894,0.0000015104034,0.000002130634,0.9973653,0.00009612115,0.0023950923,0.000030333329,0.0000012785212],"about_ca_topic_score_codex":0.005501134,"about_ca_topic_score_gemma":0.0056441855,"teacher_disagreement_score":0.005501134,"about_ca_system_score_codex":0.00089400256,"about_ca_system_score_gemma":0.00062284997,"threshold_uncertainty_score":0.010938227},"labels":[],"label_agreement":null},{"id":"W4410061435","doi":"10.1145/3676151.3719379","title":"Optimization Strategies for Enhancing Resource Efficiency in Transformers &amp; Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"","keywords":"Computer science; Transformer; Engineering; Electrical engineering; Voltage","score_opus":0.015603810564222161,"score_gpt":0.2728670736106312,"score_spread":0.257263263046409,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410061435","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038482774,0.0006356744,0.9556998,0.000906344,0.000029472156,0.000060653896,0.00011822674,0.0008383303,0.0032287056],"genre_scores_gemma":[0.65133595,0.00061648333,0.34433025,0.00027387138,0.000054276177,0.00016553416,0.0002841717,0.00038124758,0.0025581757],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99946827,0.00025295888,0.00002717039,0.00007018073,0.00012229236,0.000059206584],"domain_scores_gemma":[0.9976761,0.0018800129,0.000095187635,0.0001722311,0.00012226275,0.000054286684],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017110117,0.0006885592,0.0006750617,0.00054636656,0.00036083593,0.0011329633,0.0014172184,0.00069900404,0.0023618112],"category_scores_gemma":[0.007531988,0.00043060345,0.0005098816,0.00067990413,0.0008543236,0.0029732962,0.0011168465,0.0011608038,0.00039540665],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000092323964,0.00006555652,0.0006402896,0.00012119921,0.000033499604,0.0000511976,0.00009726007,0.87154126,0.002982887,0.042283162,0.0020934965,0.07999792],"study_design_scores_gemma":[0.0000071818176,0.00001095532,0.000040213283,0.0000046744926,0.000005536542,0.000007812341,0.000014996369,0.9863289,0.0006237979,0.0125928465,0.00036077696,0.0000023444961],"about_ca_topic_score_codex":0.004216797,"about_ca_topic_score_gemma":0.010575227,"teacher_disagreement_score":0.004216797,"about_ca_system_score_codex":0.0010973043,"about_ca_system_score_gemma":0.001386114,"threshold_uncertainty_score":0.00904882},"labels":[],"label_agreement":null},{"id":"W4410082953","doi":"10.1101/2025.05.03.25325604","title":"Investigations on using Evidence-Based GraphRag Pipeline using LLM Tailored for USMLE Style Questions","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Lakehead University","keywords":"Pipeline (software); Medical education; Computer science; Medicine; Medical physics; Programming language","score_opus":0.1722346415906235,"score_gpt":0.351912041709803,"score_spread":0.17967740011917951,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410082953","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.063881055,0.00090867135,0.84454834,0.0045859613,0.00022611674,0.0013537536,0.002349906,0.07364268,0.008503502],"genre_scores_gemma":[0.29861695,0.00024142083,0.692879,0.00095962686,0.000063074105,0.00023701835,0.0032755136,0.0012742882,0.0024531318],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9940838,0.0035317093,0.0003697139,0.0008554492,0.0009977055,0.00016161104],"domain_scores_gemma":[0.9702827,0.021083567,0.0010482869,0.00478214,0.0023207543,0.00048257332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01058625,0.0009060058,0.00073748827,0.0033720825,0.0008223045,0.0038089335,0.0023297532,0.0025275238,0.010672632],"category_scores_gemma":[0.04413027,0.00048255434,0.0012519724,0.001408971,0.0013136506,0.0058378354,0.003941659,0.0019982099,0.0036543459],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015251137,0.00097066874,0.012122644,0.0022593392,0.00030866076,0.0012841716,0.005206816,0.045605745,0.044684514,0.058646217,0.034792494,0.7925936],"study_design_scores_gemma":[0.0003842566,0.00065622997,0.0028041082,0.0004911875,0.00024734286,0.0008098475,0.0010883128,0.7832171,0.06130337,0.07126001,0.07755382,0.00018447502],"about_ca_topic_score_codex":0.0050565833,"about_ca_topic_score_gemma":0.0061082146,"teacher_disagreement_score":0.010672632,"about_ca_system_score_codex":0.0014497344,"about_ca_system_score_gemma":0.0021301361,"threshold_uncertainty_score":0.055986106},"labels":[],"label_agreement":null},{"id":"W4410170384","doi":"10.1061/jccee5.cpeng-6037","title":"The Framework and Implementation of Using Large Language Models to Answer Questions about Building Codes and Standards","year":2025,"lang":"en","type":"article","venue":"Journal of Computing in Civil Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Programming language; Architectural engineering; Engineering","score_opus":0.009319578127771067,"score_gpt":0.31825353233144277,"score_spread":0.3089339542036717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410170384","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004303125,0.00015530769,0.9520817,0.0012506399,0.000098252116,0.0004817877,0.0008508459,0.03760537,0.0031729557],"genre_scores_gemma":[0.08633276,0.00021907795,0.9041416,0.00047174754,0.00007561996,0.0008106476,0.0027628627,0.0012809824,0.0039045468],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961916,0.0019742027,0.00035051513,0.0006543072,0.0006998939,0.00012930979],"domain_scores_gemma":[0.99071115,0.0058353064,0.0005426244,0.0014866937,0.0010504472,0.00037373262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064935456,0.0012669855,0.0007066547,0.0022042277,0.0013436831,0.003655193,0.0022990361,0.0024552997,0.006433346],"category_scores_gemma":[0.018311732,0.0010346462,0.0016155059,0.0010531769,0.0014335035,0.00655485,0.0037863404,0.0033816355,0.004565234],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007635915,0.0011753656,0.007597352,0.0026398064,0.00030317152,0.0012807467,0.01283733,0.041144647,0.055489127,0.23085974,0.07431785,0.57159126],"study_design_scores_gemma":[0.00010475919,0.0002247473,0.0013001652,0.00040803736,0.00014276532,0.0005757436,0.0018214544,0.6704506,0.028602691,0.13579834,0.16033149,0.00023916789],"about_ca_topic_score_codex":0.008130498,"about_ca_topic_score_gemma":0.012997769,"teacher_disagreement_score":0.008130498,"about_ca_system_score_codex":0.0018774236,"about_ca_system_score_gemma":0.003028942,"threshold_uncertainty_score":0.034341514},"labels":[],"label_agreement":null},{"id":"W4410203956","doi":"10.1371/journal.pdig.0000800","title":"Clinical insights: A comprehensive review of language models in medicine","year":2025,"lang":"en","type":"review","venue":"PLOS Digital Health","topic":"Topic Modeling","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Saint Mary's University","funders":"","keywords":"Computer science; Context (archaeology); Key (lock); Autonomy; Data science; Resource (disambiguation); Health care; Management science; Knowledge management; Artificial intelligence; Engineering","score_opus":0.20573235257994055,"score_gpt":0.45021912129969127,"score_spread":0.24448676871975072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410203956","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000047819027,0.9979639,0.00059789093,0.0008157094,0.00010539793,0.0000070010537,0.000024026394,0.000008368464,0.00042993558],"genre_scores_gemma":[0.00111086,0.9973007,0.00077704445,0.00039400312,0.0002534748,0.000017620514,0.000036857316,0.0000053379486,0.000104145656],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983192,0.0008253616,0.00028109338,0.00018003551,0.0003444432,0.00004979649],"domain_scores_gemma":[0.98123175,0.016929677,0.00049792026,0.00024588863,0.00094395934,0.0001508455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058829593,0.0015066492,0.0025739793,0.0056006224,0.0004500315,0.0026086643,0.0018537191,0.0020899782,0.004777267],"category_scores_gemma":[0.015891053,0.00076604163,0.0015888725,0.005008433,0.0014518946,0.003612979,0.0018247145,0.002752115,0.0017613204],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006609392,0.000054516386,0.0003690505,0.045114364,0.00032880128,0.0000959736,0.0002760761,0.00094163214,0.00020166073,0.011775057,0.019435456,0.9213414],"study_design_scores_gemma":[0.000036885554,0.00015807082,0.0016695654,0.07073261,0.0009310521,0.0010118128,0.00034996925,0.0009802385,0.00037794013,0.021232339,0.9024128,0.00010674676],"about_ca_topic_score_codex":0.004063594,"about_ca_topic_score_gemma":0.005112851,"teacher_disagreement_score":0.0058829593,"about_ca_system_score_codex":0.0015597949,"about_ca_system_score_gemma":0.0045013404,"threshold_uncertainty_score":0.031112432},"labels":[],"label_agreement":null},{"id":"W4410240055","doi":"10.3390/make7020042","title":"Leveraging Failure Modes and Effect Analysis for Technical Language Processing","year":2025,"lang":"en","type":"article","venue":"Machine Learning and Knowledge Extraction","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hydro-Québec; Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada; Hydro-Québec; Université du Québec à Trois-Rivières","keywords":"Computer science","score_opus":0.008437007984600979,"score_gpt":0.30890788592712254,"score_spread":0.3004708779425216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410240055","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014141909,0.00044937318,0.97552574,0.00041804477,0.000061072686,0.00013231148,0.0013975297,0.00435074,0.0035233381],"genre_scores_gemma":[0.33495963,0.0008992415,0.65339696,0.00021995875,0.00020394879,0.0003699585,0.0060134735,0.0008684467,0.003068414],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99736696,0.00082152133,0.00025749634,0.00084793795,0.0006153192,0.000090760885],"domain_scores_gemma":[0.9869251,0.008828377,0.0010596054,0.0015924637,0.0015059432,0.000088487206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034380255,0.0014273082,0.0006332989,0.008875934,0.0007670292,0.001987778,0.00132872,0.0009781092,0.00285937],"category_scores_gemma":[0.016852869,0.00043749603,0.0017207668,0.0028630637,0.00089656806,0.0036916146,0.0018956728,0.0013530306,0.0021175817],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014749347,0.00019743302,0.014249036,0.0014317005,0.00028882475,0.0011995974,0.0030181424,0.09176915,0.03419409,0.028256424,0.0115462,0.81370187],"study_design_scores_gemma":[0.000026461668,0.00014575307,0.012674768,0.00034581823,0.00029867628,0.0009364537,0.00095349445,0.7881492,0.041860428,0.09425662,0.060172826,0.0001796069],"about_ca_topic_score_codex":0.0045136437,"about_ca_topic_score_gemma":0.0056321747,"teacher_disagreement_score":0.008875934,"about_ca_system_score_codex":0.0011716305,"about_ca_system_score_gemma":0.0014669446,"threshold_uncertainty_score":0.018182218},"labels":[],"label_agreement":null},{"id":"W4410252907","doi":"10.1016/j.procs.2025.04.497","title":"Developing Natural Language Processing Algorithms to Fact-Check Speech or Text","year":2025,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Canada West","funders":"","keywords":"Computer science; Natural language processing; Natural language; Speech recognition; Artificial intelligence; Algorithm","score_opus":0.02459784536166521,"score_gpt":0.31006999721146755,"score_spread":0.28547215184980235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410252907","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002627632,0.00012283829,0.99359834,0.0002659166,0.000038955906,0.00011684369,0.00024140993,0.0022449838,0.0007431752],"genre_scores_gemma":[0.045179565,0.0002665856,0.9515417,0.00019465652,0.00008150617,0.00015600896,0.001440409,0.00033023732,0.00080932333],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954039,0.0019668934,0.00044670628,0.0011236268,0.00092412747,0.000134674],"domain_scores_gemma":[0.9519889,0.035084367,0.002690044,0.0050221616,0.004925132,0.00028939155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00891512,0.001221268,0.00094388763,0.005463667,0.0013186485,0.004044492,0.0024150845,0.0017134765,0.0055124257],"category_scores_gemma":[0.038808268,0.00065374543,0.0016120586,0.0023959272,0.0016412109,0.007980426,0.0021869463,0.002463266,0.0045637735],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017412483,0.00028148846,0.006578485,0.001036168,0.00026099576,0.00032198537,0.0021365199,0.059269175,0.021555556,0.09080389,0.014779139,0.80280244],"study_design_scores_gemma":[0.00003405627,0.00011431135,0.0016276966,0.0002650861,0.00013251853,0.00041872094,0.0009367613,0.7585511,0.030337997,0.15790647,0.04956986,0.00010535617],"about_ca_topic_score_codex":0.003430546,"about_ca_topic_score_gemma":0.004720746,"teacher_disagreement_score":0.00891512,"about_ca_system_score_codex":0.0011554968,"about_ca_system_score_gemma":0.00220754,"threshold_uncertainty_score":0.047148287},"labels":[],"label_agreement":null},{"id":"W4410397730","doi":"10.32473/flairs.38.1.138888","title":"Leveraging Faithfulness in Abstractive Text Summarization with Elementary Discourse Units","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Natural language processing; Computer science; Linguistics; Artificial intelligence; Psychology; Philosophy","score_opus":0.11124415059323012,"score_gpt":0.3668951031233309,"score_spread":0.2556509525301008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410397730","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.109316826,0.004892684,0.8627636,0.00095620134,0.0004029653,0.00045817636,0.0025659178,0.013015579,0.005628157],"genre_scores_gemma":[0.47879457,0.0017057632,0.49917072,0.00038199828,0.00055058766,0.00031262942,0.0101008145,0.0008231969,0.008159783],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99903476,0.0002620749,0.00011776421,0.00030805962,0.00021807793,0.00005929406],"domain_scores_gemma":[0.9958474,0.001851171,0.00057390204,0.00065448106,0.0009474719,0.00012568869],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001358114,0.001259771,0.00071517687,0.0028661897,0.0005008178,0.0017432037,0.0008511275,0.0007390509,0.0021171807],"category_scores_gemma":[0.008423135,0.00028965974,0.0006451962,0.0014636194,0.0005046971,0.0030467412,0.001515771,0.0012702981,0.0017666027],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005153089,0.00017649353,0.0020490224,0.00077106524,0.00014159437,0.00019578524,0.0012639273,0.010140155,0.06175437,0.002840711,0.009431815,0.9107197],"study_design_scores_gemma":[0.00033378904,0.001588437,0.0131442705,0.00043308572,0.0009701766,0.0007229311,0.0026389593,0.628181,0.22943516,0.038716957,0.08356969,0.00026567985],"about_ca_topic_score_codex":0.0015204772,"about_ca_topic_score_gemma":0.003515931,"teacher_disagreement_score":0.0028661897,"about_ca_system_score_codex":0.00042821825,"about_ca_system_score_gemma":0.0006776933,"threshold_uncertainty_score":0.0071824193},"labels":[],"label_agreement":null},{"id":"W4410523294","doi":"10.1101/2025.05.16.652427","title":"Assessing Large Language Model Alignment Towards Radiological Myths and Misconceptions","year":2025,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Nuclear Laboratories","funders":"Atomic Energy of Canada Limited","keywords":"Mythology; Computer science; Linguistics; Epistemology; Cognitive science; Psychology; History; Philosophy; Classics","score_opus":0.027945218374378557,"score_gpt":0.27488429834201167,"score_spread":0.24693907996763312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410523294","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91284806,0.00026412366,0.0701874,0.0010703886,0.000115019466,0.00096161175,0.004899244,0.0014152123,0.0082389815],"genre_scores_gemma":[0.9420997,0.0000910245,0.046885796,0.00021553852,0.000026051412,0.001022631,0.008520228,0.00015703852,0.0009820685],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99051195,0.006337063,0.0007006777,0.0009815756,0.0012473478,0.00022139857],"domain_scores_gemma":[0.8349245,0.14659761,0.005860301,0.0043765064,0.0074248454,0.00081625837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014623274,0.0009517885,0.00048380133,0.0022711104,0.0010538313,0.0037058457,0.0008619645,0.0011459214,0.0037285662],"category_scores_gemma":[0.08194915,0.00038643347,0.0010745365,0.0016876551,0.0010300576,0.002768941,0.002530253,0.0021303592,0.0012326004],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004274474,0.0024047818,0.43303955,0.00354994,0.0013437973,0.0011829691,0.054972425,0.13587429,0.024075665,0.020848142,0.020843172,0.29759085],"study_design_scores_gemma":[0.00032010663,0.0013261548,0.10121907,0.00067608285,0.00066198205,0.0005016867,0.019475216,0.82207996,0.011800646,0.022222314,0.019391328,0.00032545],"about_ca_topic_score_codex":0.009101235,"about_ca_topic_score_gemma":0.0090181865,"teacher_disagreement_score":0.014623274,"about_ca_system_score_codex":0.0021427763,"about_ca_system_score_gemma":0.001999337,"threshold_uncertainty_score":0.07733619},"labels":[],"label_agreement":null},{"id":"W4410543982","doi":"10.14778/3717755.3717766","title":"WeShap: Weak Supervision Source Evaluation with Shapley Values","year":2024,"lang":"en","type":"article","venue":"Proceedings of the VLDB Endowment","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Shapley value; Mathematical economics; Mathematics; Game theory","score_opus":0.024034614301506736,"score_gpt":0.2549961208786427,"score_spread":0.23096150657713596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410543982","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047082994,0.00075012574,0.9351886,0.0011284319,0.00019938493,0.00035102994,0.0009115106,0.0069814003,0.0074065896],"genre_scores_gemma":[0.48160687,0.00030227096,0.5088531,0.00066363555,0.00013837783,0.0005412253,0.0024981296,0.0014511903,0.003945245],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9942194,0.0025120126,0.00033497094,0.0010508416,0.0015224309,0.000360385],"domain_scores_gemma":[0.98173565,0.011745287,0.0007999783,0.0030530875,0.0020879046,0.0005780731],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011865559,0.002104456,0.0016053612,0.002353991,0.0016765097,0.003421476,0.0036145824,0.0026567224,0.007847093],"category_scores_gemma":[0.03918682,0.00091187406,0.001352123,0.0018880274,0.002541406,0.006261868,0.0053492985,0.0043949,0.0017183601],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00096948654,0.0004139285,0.007216698,0.00068436365,0.00026391353,0.00023934181,0.00041873867,0.48785755,0.0055123325,0.0876551,0.023661265,0.3851073],"study_design_scores_gemma":[0.000055692955,0.00010963033,0.00031248873,0.000050085026,0.000021273734,0.000042201853,0.000048944552,0.9111509,0.0039689024,0.08223721,0.0019825285,0.000020133331],"about_ca_topic_score_codex":0.0033069174,"about_ca_topic_score_gemma":0.00541851,"teacher_disagreement_score":0.011865559,"about_ca_system_score_codex":0.0030443277,"about_ca_system_score_gemma":0.0052707456,"threshold_uncertainty_score":0.06275183},"labels":[],"label_agreement":null},{"id":"W4410583873","doi":"10.23919/date64628.2025.10992746","title":"Slipstream: Semantic-Based Training Acceleration for Recommendation Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; University of British Columbia","keywords":"Computer science; Acceleration; Training (meteorology); Artificial intelligence; Geography","score_opus":0.1102671493209971,"score_gpt":0.3150985974282873,"score_spread":0.20483144810729018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410583873","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.050399214,0.0014233772,0.71148384,0.00059661915,0.0005551868,0.00030083896,0.0046836836,0.2234649,0.0070922817],"genre_scores_gemma":[0.24120942,0.0006260975,0.7273135,0.0006612239,0.00015672394,0.0005436979,0.015646972,0.003826197,0.0100161135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993604,0.00008597918,0.000039306706,0.00019095263,0.00024367122,0.00007975386],"domain_scores_gemma":[0.998722,0.000462723,0.00006238437,0.00036721813,0.00029656015,0.000089023386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001033748,0.0019133012,0.00096147525,0.0010705094,0.00044080731,0.0010804608,0.0031440314,0.0011013178,0.016214833],"category_scores_gemma":[0.005984598,0.00088708074,0.0009899277,0.0015564612,0.0003754717,0.0029643113,0.001405298,0.0022062063,0.0074099055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000907315,0.0005866382,0.0052810935,0.0003719308,0.00020720116,0.0002036302,0.0001909211,0.12849674,0.011150697,0.0039538858,0.112195216,0.7364548],"study_design_scores_gemma":[0.000089304405,0.000087020875,0.000452308,0.00001177995,0.000015811012,0.000037981055,0.0000308616,0.98779744,0.0042492454,0.0020502692,0.005163572,0.000014364273],"about_ca_topic_score_codex":0.01998456,"about_ca_topic_score_gemma":0.04496682,"teacher_disagreement_score":0.01998456,"about_ca_system_score_codex":0.0010398825,"about_ca_system_score_gemma":0.002002303,"threshold_uncertainty_score":0.05424398},"labels":[],"label_agreement":null},{"id":"W4410584011","doi":"10.23919/date64628.2025.10993215","title":"Lookup Table Refactoring: Towards Efficient Logarithmic Number System Addition for Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Code refactoring; Table (database); Computer science; Logarithm; Lookup table; Programming language; Arithmetic; Parallel computing; Mathematics; Software; Database","score_opus":0.018778905387524577,"score_gpt":0.27190073522180186,"score_spread":0.2531218298342773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410584011","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04404963,0.0006409018,0.94005764,0.00028358947,0.00011843525,0.000089418885,0.00029500283,0.011158398,0.0033070636],"genre_scores_gemma":[0.36659807,0.00037660287,0.62627906,0.0002832786,0.00006951388,0.00010886724,0.0007213877,0.0010226624,0.0045405524],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993814,0.00013134626,0.00006392075,0.0001292086,0.00023929414,0.000054731077],"domain_scores_gemma":[0.99888283,0.00043122133,0.00010331728,0.00033483803,0.00021890966,0.000028859047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00058774016,0.0008092063,0.00053546834,0.00068001624,0.00033762597,0.0011870238,0.0012346138,0.0003760463,0.0050435355],"category_scores_gemma":[0.003737758,0.00031011103,0.0007808242,0.0008153403,0.00054749235,0.0021568236,0.00084759685,0.000771108,0.0020633521],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044820292,0.00012121808,0.002700298,0.0006188614,0.00007352361,0.0007767111,0.0006026498,0.12571447,0.08299481,0.0640711,0.01564327,0.7062348],"study_design_scores_gemma":[0.00006405602,0.00016696398,0.0003131844,0.00006514982,0.00006466014,0.00042501412,0.00014611032,0.8710963,0.07456285,0.032012597,0.021027727,0.000055425815],"about_ca_topic_score_codex":0.0034592755,"about_ca_topic_score_gemma":0.006627346,"teacher_disagreement_score":0.0050435355,"about_ca_system_score_codex":0.0007396283,"about_ca_system_score_gemma":0.0014477314,"threshold_uncertainty_score":0.016872287},"labels":[],"label_agreement":null},{"id":"W4410595593","doi":"10.1016/j.mlwa.2025.100666","title":"A novel unsupervised fine-tuning method for text summarization, and highlighting the limitations of ROUGE score","year":2025,"lang":"en","type":"article","venue":"Machine Learning with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Automatic summarization; ROUGE; Artificial intelligence; Computer science; Natural language processing; Machine learning","score_opus":0.03293311905183436,"score_gpt":0.27629384643799787,"score_spread":0.24336072738616352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410595593","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013958051,0.00094878214,0.97011614,0.00017877434,0.00017977518,0.000262541,0.00053184107,0.012028759,0.001795277],"genre_scores_gemma":[0.17115377,0.00038215128,0.81473505,0.00025078806,0.00029889002,0.0007602969,0.0049217506,0.0018031081,0.0056941705],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99651223,0.001042718,0.0002949238,0.0011073967,0.0008834949,0.00015918902],"domain_scores_gemma":[0.9948881,0.00140135,0.00044192138,0.00087553647,0.002226487,0.00016666719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030347896,0.0017134218,0.0011912249,0.0039132973,0.0011232895,0.0018400737,0.0013390433,0.0012453337,0.002300305],"category_scores_gemma":[0.013014719,0.0004404939,0.0011696456,0.0023860887,0.0007339494,0.0024939445,0.00134559,0.0015829229,0.002836979],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026305154,0.00018665817,0.002496505,0.00038157366,0.00025400327,0.00010189965,0.00049476215,0.028723853,0.037780497,0.0037269397,0.020387499,0.90520275],"study_design_scores_gemma":[0.00014435177,0.0008705559,0.006433936,0.0001299573,0.0002466977,0.00050816906,0.0004076938,0.87333465,0.061464604,0.014854468,0.041423503,0.00018146668],"about_ca_topic_score_codex":0.003590288,"about_ca_topic_score_gemma":0.008788537,"teacher_disagreement_score":0.0039132973,"about_ca_system_score_codex":0.00084742846,"about_ca_system_score_gemma":0.0014358438,"threshold_uncertainty_score":0.016049683},"labels":[],"label_agreement":null},{"id":"W4410602514","doi":"10.1007/s10664-025-10665-7","title":"Correction to: Utilization of pre-trained language models for adapter-based knowledge transfer in software engineering","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"","keywords":"Adapter (computing); Computer science; Software engineering; Natural language processing; Artificial intelligence; Operating system","score_opus":0.02983421994874074,"score_gpt":0.293690710703116,"score_spread":0.26385649075437523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410602514","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025932372,0.0005540893,0.0073022214,0.04275146,0.93332475,0.00010447198,0.0066899573,0.004041347,0.002638486],"genre_scores_gemma":[0.2484927,0.004996076,0.06995574,0.06627736,0.23312363,0.0015199756,0.024647566,0.012827752,0.33815917],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99201524,0.000912549,0.0022166916,0.0016283396,0.0024828857,0.000744204],"domain_scores_gemma":[0.83854115,0.041768536,0.006720295,0.017843913,0.09109244,0.004033608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044439915,0.002221977,0.0032845235,0.0044436646,0.0037106676,0.0041010003,0.0047395104,0.006869931,0.13621053],"category_scores_gemma":[0.16596168,0.0011678532,0.0016677767,0.004995053,0.002255776,0.0039021035,0.0040530884,0.00869756,0.068383776],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025017638,0.000049468385,0.00057598675,0.0004065355,0.00007019706,0.0008493999,0.00028366517,0.00028425854,0.000591884,0.0021742429,0.96268547,0.031778682],"study_design_scores_gemma":[0.00027189223,0.000119829725,0.010332727,0.0006460767,0.00012637473,0.0032144063,0.000939454,0.005710281,0.0067967484,0.010606909,0.9609492,0.00028606295],"about_ca_topic_score_codex":0.009948956,"about_ca_topic_score_gemma":0.012817865,"teacher_disagreement_score":0.13621053,"about_ca_system_score_codex":0.0036152985,"about_ca_system_score_gemma":0.004129287,"threshold_uncertainty_score":0.45566964},"labels":[],"label_agreement":null},{"id":"W4410636510","doi":"10.1145/3701716.3717531","title":"Beyond Retrieval: Generating Narratives in Conversational Recommender Systems","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Recommender system; Computer science; Narrative; Information retrieval; Natural language processing; World Wide Web; Artificial intelligence; Human–computer interaction; Linguistics","score_opus":0.02869610046458668,"score_gpt":0.2743105291464985,"score_spread":0.24561442868191183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410636510","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05812885,0.0024002532,0.92268974,0.0045128716,0.00024049434,0.0006598358,0.0019118728,0.0023525814,0.0071034315],"genre_scores_gemma":[0.47111592,0.0011848243,0.5156565,0.00068555336,0.00030371864,0.0005874699,0.004011997,0.0004733168,0.0059807203],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99411184,0.0040423996,0.00026427754,0.00083663507,0.0005564235,0.00018841155],"domain_scores_gemma":[0.9748045,0.020160658,0.0010241034,0.002120822,0.0013348026,0.00055503385],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071062557,0.0013009594,0.00089987845,0.0014934544,0.0017461061,0.0038894133,0.002156764,0.0022594936,0.004241102],"category_scores_gemma":[0.041559745,0.0009053067,0.0011920484,0.0012238404,0.001149835,0.008032557,0.0037937455,0.0025153332,0.002727368],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001531912,0.0007415172,0.021080473,0.0020382565,0.0004647555,0.0017261364,0.029397959,0.12498065,0.015588871,0.1608568,0.039598543,0.60199416],"study_design_scores_gemma":[0.00012527769,0.00019647728,0.0013480828,0.00024519223,0.00014569101,0.00047821976,0.0033824237,0.8052732,0.00961298,0.13689154,0.04217517,0.00012571344],"about_ca_topic_score_codex":0.004917223,"about_ca_topic_score_gemma":0.0064693964,"teacher_disagreement_score":0.0071062557,"about_ca_system_score_codex":0.0011194235,"about_ca_system_score_gemma":0.0015803191,"threshold_uncertainty_score":0.03758192},"labels":[],"label_agreement":null},{"id":"W4410673605","doi":"10.1075/term.00082.tra","title":"LlamATE","year":2025,"lang":"en","type":"article","venue":"Terminology International Journal of Theoretical and Applied Issues in Specialized Communication","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science","score_opus":0.012378580897841296,"score_gpt":0.3148047650703796,"score_spread":0.3024261841725383,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410673605","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029863372,0.0030503797,0.16340296,0.004052841,0.002665109,0.0011292732,0.22392075,0.37530106,0.19661418],"genre_scores_gemma":[0.13652572,0.0011418876,0.13173828,0.0028487511,0.00059665705,0.001689711,0.54012185,0.0318852,0.153452],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99707425,0.00073260634,0.00018479097,0.0009409256,0.000794262,0.00027317318],"domain_scores_gemma":[0.9951964,0.0011094152,0.00020201263,0.0019404904,0.0012147212,0.0003369066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025793617,0.0017826536,0.0012178505,0.0021919298,0.0012434603,0.003939084,0.0028533076,0.0019265779,0.12696505],"category_scores_gemma":[0.011028591,0.000732703,0.0015748806,0.0016414798,0.0006727496,0.004354606,0.005223556,0.0021446673,0.14122853],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009330785,0.00025704395,0.0060150786,0.0013764465,0.00012320478,0.0003751345,0.0005368134,0.008143712,0.0052952333,0.014592136,0.6569087,0.3054434],"study_design_scores_gemma":[0.00013890972,0.00022289393,0.0035780834,0.0002302834,0.000046722085,0.00040436274,0.00025024358,0.031531546,0.006434889,0.014836804,0.94223034,0.00009497497],"about_ca_topic_score_codex":0.004839198,"about_ca_topic_score_gemma":0.0059365625,"teacher_disagreement_score":0.12696505,"about_ca_system_score_codex":0.0010464231,"about_ca_system_score_gemma":0.0017087205,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4410785635","doi":"10.3389/frai.2025.1592399","title":"Moving LLM evaluation forward: lessons from human judgment research","year":2025,"lang":"en","type":"article","venue":"Frontiers in Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Quebec - Clinical Research Organization in Cancer","funders":"","keywords":"Parallels; Management science; Engineering ethics; Computer science; Psychology; Epistemology; Sociology; Engineering","score_opus":0.1868207403157975,"score_gpt":0.4347521440215913,"score_spread":0.24793140370579378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410785635","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07365201,0.024867445,0.7448116,0.0996585,0.0015413987,0.0005498113,0.00047386484,0.0015242286,0.052921154],"genre_scores_gemma":[0.7317971,0.0048359283,0.25379512,0.005314744,0.00095741515,0.00035910925,0.0003562477,0.00057941244,0.0020048502],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.861655,0.10529573,0.0047361176,0.006291251,0.02072189,0.0013000609],"domain_scores_gemma":[0.56283915,0.35309452,0.011906315,0.028043134,0.039567377,0.00454955],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14287253,0.0016913925,0.0020686768,0.0069245636,0.002395486,0.015546924,0.004619963,0.0041648103,0.006021155],"category_scores_gemma":[0.45144975,0.0007202115,0.0010784052,0.003894262,0.011312458,0.03141399,0.00900379,0.007183075,0.0016347738],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059449906,0.0005121577,0.018497113,0.0023411731,0.00048612122,0.00022482422,0.022700299,0.01709274,0.0014789052,0.30578128,0.021619285,0.60867155],"study_design_scores_gemma":[0.000097375734,0.0002629727,0.0045028967,0.001666318,0.000079777004,0.00013213264,0.006733189,0.044825107,0.0019710297,0.908754,0.030759778,0.00021542258],"about_ca_topic_score_codex":0.009757768,"about_ca_topic_score_gemma":0.008562386,"teacher_disagreement_score":0.8571275,"about_ca_system_score_codex":0.006075396,"about_ca_system_score_gemma":0.0068677417,"threshold_uncertainty_score":0.75559115},"labels":[],"label_agreement":null},{"id":"W4410794822","doi":"10.1177/08944393251344865","title":"Prompting the Machine: Introducing an LLM Data Extraction Method for Social Scientists","year":2025,"lang":"en","type":"article","venue":"Social Science Computer Review","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal; Université Laval","funders":"","keywords":"Computer science; Extraction (chemistry); Data extraction; Data science; Political science; MEDLINE; Chromatography; Chemistry","score_opus":0.10567370972666913,"score_gpt":0.46011694349038995,"score_spread":0.35444323376372083,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410794822","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013804125,0.00012584653,0.98904043,0.001985813,0.00016296195,0.001954462,0.000906766,0.0028033673,0.001639934],"genre_scores_gemma":[0.0052208146,0.0000968111,0.9894695,0.00033399972,0.000060995026,0.0032884087,0.00052770466,0.00051273673,0.00048901327],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.8845164,0.08762831,0.010870065,0.004789598,0.011397224,0.00079835724],"domain_scores_gemma":[0.7292083,0.18645796,0.012655985,0.0463017,0.023971563,0.0014045369],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1282117,0.001391391,0.0012440657,0.009563798,0.00291445,0.010629243,0.0042213807,0.0026956354,0.008368352],"category_scores_gemma":[0.28166494,0.001932982,0.0025811866,0.0071004312,0.0038909933,0.011972182,0.014409798,0.00524304,0.006069045],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039988445,0.00029295846,0.005414987,0.004633668,0.00028953905,0.00053864246,0.04186461,0.003641615,0.010897332,0.22786307,0.047465533,0.65669805],"study_design_scores_gemma":[0.0002744042,0.00026879,0.0036228027,0.0044467226,0.00016843677,0.00047608628,0.00716556,0.04331813,0.017622843,0.3607984,0.56140274,0.00043505005],"about_ca_topic_score_codex":0.0021741074,"about_ca_topic_score_gemma":0.0045184987,"teacher_disagreement_score":0.87178826,"about_ca_system_score_codex":0.00300799,"about_ca_system_score_gemma":0.011774629,"threshold_uncertainty_score":0.6780564},"labels":[],"label_agreement":null},{"id":"W4410877227","doi":"10.1007/978-3-031-94575-5_7","title":"Predicting the Road Ahead: A Knowledge Graph Based Foundation Model for Scene Understanding in Autonomous Driving","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Foundation (evidence); Graph; Artificial intelligence; Theoretical computer science","score_opus":0.048865887956767474,"score_gpt":0.2792497001078019,"score_spread":0.23038381215103446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410877227","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.090017825,0.00051587704,0.90185547,0.00036426034,0.000044891247,0.00007440764,0.0021137143,0.0020575589,0.0029559827],"genre_scores_gemma":[0.85690457,0.0005279613,0.13634373,0.000093148,0.000039055707,0.00011750668,0.0032105413,0.0001666414,0.0025968573],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99983907,0.000022254595,0.000010329351,0.000064331725,0.00003745836,0.000026465517],"domain_scores_gemma":[0.99948287,0.00030627134,0.000046289162,0.00004458476,0.000086709864,0.000033125696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002741079,0.0005622566,0.0007561377,0.0013378137,0.00046029186,0.0008734106,0.0016331241,0.0009922845,0.0026422685],"category_scores_gemma":[0.001351653,0.0005308019,0.0009777448,0.0012356591,0.0003836734,0.0019255751,0.0008469191,0.0010548211,0.00062129454],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001463034,0.00013738389,0.0019657998,0.00008441695,0.000080491474,0.00009345067,0.00011001618,0.8737568,0.0020647212,0.010477833,0.0031885207,0.10789422],"study_design_scores_gemma":[0.0000026580958,0.0000075176476,0.00019064086,0.000003997915,0.000009166447,0.0000050153126,0.0000067726255,0.99419725,0.00015426146,0.0052352296,0.00018482038,0.0000028028269],"about_ca_topic_score_codex":0.04281283,"about_ca_topic_score_gemma":0.047857404,"teacher_disagreement_score":0.04281283,"about_ca_system_score_codex":0.000874349,"about_ca_system_score_gemma":0.0010972369,"threshold_uncertainty_score":0.085127234},"labels":[],"label_agreement":null},{"id":"W4411001159","doi":"10.1101/2025.06.02.657491","title":"Cortical language areas are coupled via a soft hierarchy of model-based linguistic features","year":2025,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Children's Hospital; University of British Columbia","funders":"National Institutes of Health; BC Children's Hospital","keywords":"Hierarchy; Linguistics; Computer science; Language model; Artificial intelligence; Natural language processing; Philosophy; Political science","score_opus":0.013219448127636672,"score_gpt":0.2392853403664169,"score_spread":0.22606589223878024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411001159","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88230366,0.000057290614,0.11612848,0.000108513996,0.0000037666578,0.00001446856,0.000055987497,0.0001637605,0.001163975],"genre_scores_gemma":[0.9948613,0.000017999211,0.0048614494,0.0000068384506,0.0000019852025,0.0000086766,0.000026685517,0.0000139195145,0.00020117073],"study_design_codex":"bench_or_experimental","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99980646,0.00006634034,0.000007963692,0.000067574445,0.000026750193,0.000024906],"domain_scores_gemma":[0.9995023,0.00023783959,0.0001231381,0.00005929901,0.000035329216,0.000042081192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00040399286,0.0002731096,0.0002554398,0.00039517073,0.00019265448,0.0005892425,0.0003357236,0.00030151676,0.0008918815],"category_scores_gemma":[0.0021633408,0.00031819486,0.00034779642,0.00020623735,0.00059676915,0.001033949,0.0006943263,0.00043970795,0.000084919324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004930694,0.00016401762,0.03481903,0.00018870464,0.00036597505,0.0006306824,0.0013871276,0.36486134,0.50772953,0.036208477,0.0007424737,0.052409574],"study_design_scores_gemma":[0.000016173126,0.000110680434,0.034789916,0.000009540972,0.00004886983,0.0001430025,0.00015288268,0.92226064,0.014958207,0.027179653,0.00030894345,0.00002161281],"about_ca_topic_score_codex":0.001747755,"about_ca_topic_score_gemma":0.0021032458,"teacher_disagreement_score":0.001747755,"about_ca_system_score_codex":0.0003375415,"about_ca_system_score_gemma":0.00023927151,"threshold_uncertainty_score":0.0034751296},"labels":[],"label_agreement":null},{"id":"W4411030024","doi":"10.3390/fi17060252","title":"LLM4Rec: A Comprehensive Survey on the Integration of Large Language Models in Recommender Systems—Approaches, Applications and Challenges","year":2025,"lang":"en","type":"article","venue":"Future Internet","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Recommender system; Data science; World Wide Web; Information retrieval","score_opus":0.11543423099620027,"score_gpt":0.29231144998221414,"score_spread":0.17687721898601388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411030024","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007112327,0.41368827,0.55251956,0.007145621,0.0010495846,0.00035442255,0.0015635738,0.0030016254,0.013565059],"genre_scores_gemma":[0.0883171,0.4149622,0.4713511,0.0039844993,0.003937433,0.00068783236,0.0056301937,0.0011926545,0.009936997],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946138,0.0025441523,0.00057403283,0.00069142296,0.0013958638,0.0001806848],"domain_scores_gemma":[0.97407496,0.019627739,0.00043954153,0.002605386,0.002876416,0.00037603045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01069787,0.0015686981,0.0033539187,0.0046246164,0.0009527848,0.0042578424,0.0032458524,0.0022934296,0.0061842566],"category_scores_gemma":[0.031082777,0.0015697962,0.0026361695,0.007820538,0.0006918411,0.007367736,0.0033555098,0.0034285027,0.0046434696],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009344733,0.00020936469,0.0039011452,0.0046191304,0.0005746577,0.00010760349,0.0005252674,0.013530629,0.0012375682,0.034829855,0.034388453,0.9059828],"study_design_scores_gemma":[0.00005439606,0.00055797317,0.0059694373,0.004473154,0.0009522496,0.00086088415,0.0008208827,0.24398117,0.0031467378,0.09892873,0.6399163,0.00033810554],"about_ca_topic_score_codex":0.009695108,"about_ca_topic_score_gemma":0.011206428,"teacher_disagreement_score":0.01069787,"about_ca_system_score_codex":0.0020128367,"about_ca_system_score_gemma":0.0027764617,"threshold_uncertainty_score":0.05657643},"labels":[],"label_agreement":null},{"id":"W4411116044","doi":"10.1007/s12559-025-10470-w","title":"From RNNs to Transformers and Beyond: a Deep Dive into Intent Detection in Goal-oriented Conversational Agents","year":2025,"lang":"en","type":"article","venue":"Cognitive Computation","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Transformer; Computer science; Artificial intelligence; Recurrent neural network; Artificial neural network; Engineering; Electrical engineering","score_opus":0.012638123791752153,"score_gpt":0.27388219022103344,"score_spread":0.2612440664292813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411116044","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03084968,0.0010998964,0.9620326,0.0010792346,0.00007244405,0.000052321957,0.00019654754,0.0010265614,0.0035906727],"genre_scores_gemma":[0.79414195,0.0011512106,0.20008542,0.00034604187,0.00012405484,0.00007407999,0.00043072258,0.00025438942,0.0033921713],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99929297,0.00034057492,0.000046283105,0.00016617603,0.000088180546,0.000065808665],"domain_scores_gemma":[0.9965186,0.0025920998,0.00015231928,0.0003156514,0.0002561184,0.00016520626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014886304,0.0009770115,0.000705164,0.00088877825,0.0005007371,0.0028551286,0.0014389806,0.0010475059,0.002991985],"category_scores_gemma":[0.009152885,0.0006430495,0.0008466185,0.0007157666,0.0009940005,0.00584417,0.0019690841,0.0030497117,0.00083843304],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000504043,0.00022607775,0.006480227,0.00068405195,0.0002711217,0.0003013173,0.003124487,0.1707151,0.016297922,0.23171104,0.007037048,0.5626475],"study_design_scores_gemma":[0.000006356757,0.000032915184,0.0004278267,0.000061595965,0.000028055965,0.000038208036,0.00020519756,0.7911804,0.0019838891,0.20399427,0.0020258501,0.000015497168],"about_ca_topic_score_codex":0.0076722177,"about_ca_topic_score_gemma":0.007929212,"teacher_disagreement_score":0.0076722177,"about_ca_system_score_codex":0.0009261223,"about_ca_system_score_gemma":0.0011511511,"threshold_uncertainty_score":0.015255153},"labels":[],"label_agreement":null},{"id":"W4411117117","doi":"10.18653/v1/w14-4407","title":"A Template-based Abstractive Meeting Summarization: Leveraging Summary and Source Text Relationships","year":2014,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Information retrieval; Natural language processing; Artificial intelligence; World Wide Web","score_opus":0.032808899971617525,"score_gpt":0.22848764582304318,"score_spread":0.19567874585142564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411117117","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010876898,0.00067514944,0.960574,0.00031898098,0.00024934474,0.0003846229,0.0025427877,0.022869812,0.0015083891],"genre_scores_gemma":[0.06297512,0.0005010482,0.9210466,0.00013977845,0.00032739618,0.00035443107,0.010392322,0.0011227048,0.003140536],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973214,0.0006744977,0.0003425238,0.0008584702,0.00072059716,0.00008246027],"domain_scores_gemma":[0.9929456,0.0023439424,0.001122258,0.0010414283,0.002278537,0.00026820577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002432944,0.0020228445,0.0012250929,0.0037482132,0.0008407143,0.0026434576,0.0019645856,0.0012830237,0.004501939],"category_scores_gemma":[0.011278127,0.00064357335,0.0012403275,0.0022786963,0.00034240392,0.0030200721,0.0014730892,0.0014100282,0.005987875],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048741946,0.00018301692,0.0017365817,0.0011861266,0.00024849217,0.0003467722,0.001295619,0.0079444125,0.09965194,0.0026176365,0.02889028,0.85541165],"study_design_scores_gemma":[0.00023663032,0.0011820775,0.00880562,0.0003137843,0.0011513898,0.0016320137,0.0015263301,0.5757142,0.25592762,0.0143238725,0.13875146,0.00043496466],"about_ca_topic_score_codex":0.0020054192,"about_ca_topic_score_gemma":0.0022516535,"teacher_disagreement_score":0.004501939,"about_ca_system_score_codex":0.00042631858,"about_ca_system_score_gemma":0.0010571555,"threshold_uncertainty_score":0.015060484},"labels":[],"label_agreement":null},{"id":"W4411119063","doi":"10.18653/v1/2025.nlp4dh-1.25","title":"Evaluating Large Language Models for Narrative Topic Labeling","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Computer science; Narrative; Natural language processing; Linguistics; Philosophy","score_opus":0.0665263816001478,"score_gpt":0.3842538718380301,"score_spread":0.3177274902378823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411119063","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54371166,0.005896426,0.42322847,0.002163096,0.000534638,0.0011906547,0.0039321203,0.0116290245,0.007713916],"genre_scores_gemma":[0.80109245,0.00053382246,0.18706591,0.0003528987,0.0002009278,0.00063171133,0.0072583514,0.00043557855,0.00242838],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9873039,0.009083004,0.0005825428,0.0018411563,0.0009152361,0.00027410642],"domain_scores_gemma":[0.9430758,0.049741484,0.0012537843,0.0027272678,0.0023889733,0.0008127436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020498127,0.002295603,0.0013523431,0.002970815,0.0014230785,0.0028715625,0.0022953595,0.0022214022,0.002305634],"category_scores_gemma":[0.049035445,0.0007775653,0.0012453847,0.001595454,0.00094008667,0.003954964,0.0023189124,0.0021950179,0.0017262807],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003522445,0.0014587694,0.026879517,0.0014150923,0.0013387619,0.00036913762,0.0022292994,0.42016873,0.011419761,0.005584759,0.014908177,0.5107056],"study_design_scores_gemma":[0.00009964107,0.0002333441,0.0016637915,0.000045067536,0.00007756399,0.00005279185,0.0004075357,0.98802245,0.003906856,0.0036980112,0.0017589993,0.00003383382],"about_ca_topic_score_codex":0.010891061,"about_ca_topic_score_gemma":0.018225702,"teacher_disagreement_score":0.020498127,"about_ca_system_score_codex":0.0026438874,"about_ca_system_score_gemma":0.00170714,"threshold_uncertainty_score":0.10840577},"labels":[],"label_agreement":null},{"id":"W4411119199","doi":"10.18653/v1/2025.wnu-1.1","title":"NarraDetect: An annotated dataset for the task of narrative detection","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Computer science; Task (project management); Narrative; Natural language processing; Artificial intelligence; Information retrieval; Engineering; Art; Literature","score_opus":0.022998595146991565,"score_gpt":0.29581611807991254,"score_spread":0.272817522932921,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411119199","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029292947,0.0019267837,0.0068713706,0.0007417557,0.00047570092,0.00037478667,0.9410764,0.008106065,0.011134239],"genre_scores_gemma":[0.0155458385,0.00027285772,0.011784048,0.00013483407,0.00006108512,0.00042384255,0.9677528,0.00030730633,0.0037173168],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987784,0.00032324973,0.00017300266,0.00038469973,0.00025994438,0.00008072554],"domain_scores_gemma":[0.996376,0.0015770578,0.000359557,0.0007548179,0.00062932365,0.00030332513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009771422,0.0017083993,0.00063167163,0.0037306836,0.0014432067,0.0013823971,0.0019619528,0.002462467,0.012083024],"category_scores_gemma":[0.006670268,0.00038317707,0.0009340148,0.0032835184,0.0006782752,0.0023063791,0.0018746662,0.0016953504,0.01198208],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045966846,0.00029690829,0.005465348,0.0035960455,0.00012122849,0.001229349,0.0017699597,0.0020263747,0.0080672065,0.0029247538,0.9213638,0.052679334],"study_design_scores_gemma":[0.00020972855,0.00015581452,0.024848824,0.000560151,0.000087617336,0.0016485922,0.0019179217,0.011531881,0.0069481772,0.0027077901,0.9492416,0.00014197035],"about_ca_topic_score_codex":0.012889397,"about_ca_topic_score_gemma":0.04591913,"teacher_disagreement_score":0.012889397,"about_ca_system_score_codex":0.0012762819,"about_ca_system_score_gemma":0.0013939051,"threshold_uncertainty_score":0.040421724},"labels":[],"label_agreement":null},{"id":"W4411119578","doi":"10.18653/v1/2025.knowledgenlp-1.1","title":"Entity Retrieval for Answering Entity-Centric Questions","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Question answering; Computer science; Information retrieval; Natural language processing; Artificial intelligence; World Wide Web","score_opus":0.016701432455897558,"score_gpt":0.2788969821843046,"score_spread":0.26219554972840703,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411119578","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02541365,0.0023780032,0.96701217,0.00037751417,0.000057857796,0.00031643,0.00067448564,0.0015456743,0.0022242009],"genre_scores_gemma":[0.37368146,0.0021387835,0.61687475,0.00024383535,0.00023464204,0.00034600822,0.0034616124,0.00013158072,0.0028873289],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99758077,0.0012461026,0.00016591282,0.00046282302,0.00044328012,0.00010099723],"domain_scores_gemma":[0.9946964,0.0034241634,0.00037558816,0.0008516672,0.00056818814,0.00008397056],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031727294,0.0007843217,0.00087068306,0.0030129333,0.0006585854,0.0014857415,0.0012227678,0.0014349252,0.0038253793],"category_scores_gemma":[0.012388194,0.00028856078,0.0008492768,0.0029609124,0.0006400458,0.004717093,0.0014470394,0.001032654,0.0018211043],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042878572,0.00045698747,0.004832127,0.0017858925,0.00023326508,0.00025335175,0.0014208887,0.04884514,0.058409426,0.06335994,0.01473504,0.8052391],"study_design_scores_gemma":[0.00009195904,0.00034796415,0.0045348536,0.000106262814,0.00021902648,0.0008158296,0.0005321409,0.86519456,0.03055715,0.06975922,0.027741583,0.00009945437],"about_ca_topic_score_codex":0.0025413744,"about_ca_topic_score_gemma":0.0033021853,"teacher_disagreement_score":0.0038253793,"about_ca_system_score_codex":0.0005616511,"about_ca_system_score_gemma":0.0007382196,"threshold_uncertainty_score":0.016779184},"labels":[],"label_agreement":null},{"id":"W4411120250","doi":"10.18653/v1/2025.wnu-1.7","title":"A Theoretical Framework for Evaluating Narrative Surprise in Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Surprise; Narrative; Computer science; Natural language processing; Linguistics; Psychology; Philosophy; Communication","score_opus":0.029949040192777384,"score_gpt":0.3667034877914282,"score_spread":0.3367544475986508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411120250","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03394982,0.00011826567,0.9623942,0.0004848749,0.000015532205,0.00023587391,0.00019017096,0.000285764,0.0023254715],"genre_scores_gemma":[0.64012104,0.00011555927,0.35703167,0.0001662832,0.0000687993,0.0015066784,0.00046937546,0.00008415081,0.0004363079],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9832659,0.011695566,0.00088055816,0.001593603,0.0021625182,0.0004019649],"domain_scores_gemma":[0.9018181,0.08150781,0.0076365834,0.004329075,0.003643264,0.0010651615],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0308545,0.0018851829,0.0011395405,0.0045075496,0.0014932111,0.0041652573,0.0025295252,0.0016735411,0.0031967636],"category_scores_gemma":[0.12951958,0.0008244696,0.0016009712,0.002513391,0.0042690695,0.006028774,0.004415776,0.0023427443,0.0004190271],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055768585,0.00058602024,0.03820621,0.0008557733,0.00055301987,0.00044963026,0.007013972,0.28070018,0.009596571,0.5540126,0.0023299106,0.105138496],"study_design_scores_gemma":[0.000053887306,0.00043140916,0.0037283513,0.000079203164,0.00008025587,0.00015350465,0.0010065684,0.6653764,0.0020527837,0.32466084,0.0023044327,0.000072301715],"about_ca_topic_score_codex":0.002114374,"about_ca_topic_score_gemma":0.0018349244,"teacher_disagreement_score":0.0308545,"about_ca_system_score_codex":0.002404009,"about_ca_system_score_gemma":0.0018543659,"threshold_uncertainty_score":0.16317618},"labels":[],"label_agreement":null},{"id":"W4411127324","doi":"10.22399/ijcesen.2481","title":"INTELLIDOC - An Adaptive Transformer-Powered Pipeline For Intelligent Document Processing And Entity Extraction","year":2025,"lang":"en","type":"article","venue":"International Journal of Computational and Experimental Science and Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Transformer; Pipeline (software); Computer science; Artificial intelligence; Natural language processing; Engineering; Electrical engineering; Programming language","score_opus":0.014874400241966297,"score_gpt":0.3138410068152804,"score_spread":0.2989666065733141,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411127324","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01979173,0.0013675388,0.8726218,0.0002840063,0.00016074476,0.00034781185,0.0030295926,0.09801447,0.004382268],"genre_scores_gemma":[0.1123231,0.00061084033,0.8574222,0.0004139389,0.00009060492,0.0003072377,0.015180453,0.0019439921,0.011707661],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992773,0.00006454735,0.00005638474,0.0002650184,0.00025223804,0.00008460655],"domain_scores_gemma":[0.99842083,0.00038660027,0.00010953172,0.0005326091,0.0004780218,0.00007245262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00088827667,0.0012670845,0.00089466164,0.0022433505,0.00057061797,0.001788095,0.002385949,0.0008481726,0.004456951],"category_scores_gemma":[0.0026965265,0.00054894446,0.0010208557,0.0021592304,0.00056720024,0.0034724707,0.001787393,0.0011378499,0.0060262573],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004631852,0.00017223078,0.0018353577,0.00043969176,0.00010708605,0.0002706409,0.00029122308,0.009104806,0.05588825,0.0032715825,0.05426002,0.873896],"study_design_scores_gemma":[0.00019043358,0.0005217843,0.0035292327,0.00008382615,0.00020429256,0.0013926043,0.00033447836,0.6308545,0.22187114,0.010364504,0.13046904,0.00018419843],"about_ca_topic_score_codex":0.007042459,"about_ca_topic_score_gemma":0.012793348,"teacher_disagreement_score":0.007042459,"about_ca_system_score_codex":0.0010383461,"about_ca_system_score_gemma":0.0020269493,"threshold_uncertainty_score":0.014909923},"labels":[],"label_agreement":null},{"id":"W4411143753","doi":"10.1109/icaace65325.2025.11019781","title":"Zero-Shot End-to-End Relation Extraction in Chinese: A Comparative Study of Gemini, LLaMA, and ChatGPT","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"End-to-end principle; Zero (linguistics); Shot (pellet); Relation (database); Extraction (chemistry); Computer science; Mathematics; Artificial intelligence; Chromatography; Materials science; Chemistry; Data mining; Linguistics","score_opus":0.044190243230717914,"score_gpt":0.3451886114867058,"score_spread":0.3009983682559879,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411143753","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.760631,0.009709532,0.16475561,0.0015647333,0.0006366456,0.00079343654,0.00996897,0.038160395,0.013779579],"genre_scores_gemma":[0.8234739,0.0019768872,0.13797191,0.0005082701,0.00012187527,0.00029386792,0.02809922,0.0011033865,0.0064507104],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9973248,0.0007887836,0.0002680739,0.00097206066,0.00043913853,0.00020703375],"domain_scores_gemma":[0.9935933,0.0040027886,0.0001784714,0.0012586561,0.0007538409,0.00021303388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050500673,0.0019137602,0.0015098549,0.002613014,0.0013220764,0.0019060939,0.0025422135,0.0010372704,0.0026736138],"category_scores_gemma":[0.014650514,0.000541238,0.001192926,0.0023199718,0.0009484428,0.0065474296,0.0024018702,0.0015223998,0.0020251097],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002895874,0.0007269258,0.029458933,0.0026516113,0.0009311574,0.0016583503,0.0032558567,0.034902688,0.027806455,0.0046725185,0.030514408,0.8605252],"study_design_scores_gemma":[0.0002937805,0.0012406663,0.045449812,0.00023957506,0.0010726447,0.002499779,0.0047631264,0.81485134,0.075879514,0.009887308,0.043444477,0.0003779054],"about_ca_topic_score_codex":0.036606673,"about_ca_topic_score_gemma":0.050742324,"teacher_disagreement_score":0.036606673,"about_ca_system_score_codex":0.0015126021,"about_ca_system_score_gemma":0.0028268301,"threshold_uncertainty_score":0.072787166},"labels":[],"label_agreement":null},{"id":"W4411170507","doi":"10.1080/00295450.2025.2481358","title":"Natural Language Processing in the Nuclear Industry: Opportunities and Challenges","year":2025,"lang":"en","type":"article","venue":"Nuclear Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Nuclear industry; Computer science; Nuclear engineering; Engineering","score_opus":0.04038123152615424,"score_gpt":0.26083551931123056,"score_spread":0.2204542877850763,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411170507","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014414505,0.53066915,0.21680738,0.20449023,0.0018383419,0.00019330412,0.00087838364,0.0007501579,0.029958488],"genre_scores_gemma":[0.13017465,0.6292981,0.2097029,0.018824788,0.004653267,0.000556409,0.0015622625,0.0003454649,0.0048821545],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9907532,0.0053551253,0.00069039536,0.0009479575,0.0020224655,0.00023085024],"domain_scores_gemma":[0.93410754,0.058731463,0.0013442931,0.0015173798,0.0037946342,0.0005046797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01757457,0.0007522825,0.0012382499,0.0031040362,0.0013332823,0.008009981,0.0026383195,0.0043648537,0.0035676043],"category_scores_gemma":[0.02343213,0.0006244405,0.00087047194,0.0052210046,0.0060178344,0.016598422,0.0034410788,0.0043260786,0.0018569573],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008034737,0.0001459312,0.002693541,0.008646039,0.000097528995,0.00067706406,0.0033212428,0.010314858,0.0023624133,0.25428814,0.028170815,0.68920213],"study_design_scores_gemma":[0.000020315716,0.000069438574,0.0017141699,0.005243509,0.00004379695,0.000784864,0.006535352,0.02170556,0.0015617581,0.5199907,0.44220433,0.000126117],"about_ca_topic_score_codex":0.004482568,"about_ca_topic_score_gemma":0.0039646025,"teacher_disagreement_score":0.01757457,"about_ca_system_score_codex":0.002634771,"about_ca_system_score_gemma":0.0059794826,"threshold_uncertainty_score":0.092944264},"labels":[],"label_agreement":null},{"id":"W4411171689","doi":"10.1109/access.2025.3578497","title":"Unsupervised Context-Linking Retriever for Question Answering on Long Narrative Books","year":2025,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Universiti Kebangsaan Malaysia; Ministry of Higher Education, Malaysia","keywords":"Question answering; Computer science; Context (archaeology); Narrative; Information retrieval; Labrador Retriever; Natural language processing; Artificial intelligence; World Wide Web; Linguistics; Medicine; History","score_opus":0.04394293215499163,"score_gpt":0.338799564082537,"score_spread":0.29485663192754535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411171689","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13548599,0.0044787186,0.79072076,0.00046575165,0.0002650815,0.0006280232,0.004100028,0.055046987,0.008808639],"genre_scores_gemma":[0.3575554,0.0013048635,0.60375655,0.00072068284,0.0003331537,0.0006026985,0.01916109,0.0015300318,0.015035529],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989813,0.00029986573,0.00006834976,0.0003864237,0.00020250591,0.00006159735],"domain_scores_gemma":[0.9984034,0.00083741447,0.00010683207,0.00026842626,0.00032659853,0.00005735305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010391047,0.0016485748,0.0010983356,0.0015280043,0.0006219264,0.0009095503,0.0017088458,0.0014632508,0.0066481754],"category_scores_gemma":[0.0055350885,0.00034748908,0.0009924627,0.0009162513,0.00040160638,0.0028316758,0.0015169566,0.0014793215,0.0074808286],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077044166,0.00045088044,0.0020687953,0.0011285595,0.00019864003,0.0011510977,0.00097582804,0.04644255,0.083863586,0.00481291,0.035365332,0.8227715],"study_design_scores_gemma":[0.00014956015,0.0006955223,0.0027377245,0.00009163482,0.00019527631,0.0010498088,0.0007280928,0.8536519,0.10044665,0.00879171,0.03135622,0.00010585728],"about_ca_topic_score_codex":0.003598335,"about_ca_topic_score_gemma":0.006679382,"teacher_disagreement_score":0.0066481754,"about_ca_system_score_codex":0.00053781434,"about_ca_system_score_gemma":0.0009401727,"threshold_uncertainty_score":0.0222404},"labels":[],"label_agreement":null},{"id":"W4411199189","doi":"10.2196/76773","title":"Toward Cross-Hospital Deployment of Natural Language Processing Systems: Model Development and Validation of Fine-Tuned Large Language Models for Disease Name Recognition in Japanese","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Software deployment; Computer science; Robustness (evolution); Artificial intelligence; Natural language processing; Medicine; World Wide Web; Software engineering","score_opus":0.02115195569370589,"score_gpt":0.30708802305490746,"score_spread":0.2859360673612016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411199189","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89586306,0.0016881203,0.09335351,0.0006953171,0.0002947397,0.00059263455,0.0015906969,0.0037875709,0.002134277],"genre_scores_gemma":[0.9405956,0.00031240835,0.052482836,0.0002640691,0.000057229765,0.00048334824,0.004325243,0.00017840991,0.0013008689],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99777406,0.00085893605,0.00020610394,0.00080378447,0.00016465201,0.00019248301],"domain_scores_gemma":[0.9933381,0.0040021245,0.0003711146,0.0006867602,0.0013233698,0.00027848492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072753713,0.0019284971,0.00092922867,0.001125512,0.0007051582,0.0012763032,0.0019393696,0.0014078928,0.001232968],"category_scores_gemma":[0.014431769,0.00075926346,0.0015288849,0.0007235297,0.00066451065,0.0019447088,0.0020982667,0.0022289783,0.000798408],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014129257,0.0012752278,0.03359226,0.00049437035,0.000831371,0.0005491719,0.0009618633,0.73829246,0.014236621,0.00097544363,0.0060254163,0.20135294],"study_design_scores_gemma":[0.000052096857,0.00020212845,0.00407044,0.00001803547,0.00009402445,0.000041754105,0.00012973914,0.9905104,0.0039272644,0.0003846358,0.000538462,0.000030971427],"about_ca_topic_score_codex":0.049327094,"about_ca_topic_score_gemma":0.036986582,"teacher_disagreement_score":0.049327094,"about_ca_system_score_codex":0.0022514295,"about_ca_system_score_gemma":0.002279799,"threshold_uncertainty_score":0.09807998},"labels":[],"label_agreement":null},{"id":"W4411267758","doi":"10.22318/icls2025.100255","title":"A Systematic Literature Review of Large Language Model Applications in the Algebra Domain","year":2025,"lang":"en","type":"article","venue":"Proceedings.","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Algebra over a field; Programming language; Domain (mathematical analysis); Mathematics; Pure mathematics","score_opus":0.00780541389418628,"score_gpt":0.27125327806474,"score_spread":0.2634478641705537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411267758","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022388883,0.9935136,0.0011692854,0.0007175294,0.00014225898,0.00042141386,0.0009510226,0.00002572738,0.0008201889],"genre_scores_gemma":[0.021065317,0.9733865,0.002828084,0.00073844875,0.00007240626,0.0010232395,0.0006959886,0.000017558465,0.00017249267],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.986053,0.005728052,0.0051779468,0.0010025246,0.0017845883,0.00025392429],"domain_scores_gemma":[0.9280941,0.05790661,0.006736353,0.0014336861,0.0052696536,0.0005595551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016369607,0.0012129789,0.0043815523,0.017252557,0.0009313449,0.0029473226,0.0018562814,0.0015235969,0.004648921],"category_scores_gemma":[0.087866634,0.0009866712,0.005054739,0.016690448,0.001282126,0.0036511905,0.0027994958,0.0013453439,0.0005894509],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000113847396,0.00001807843,0.0010590907,0.90989643,0.003648574,0.0001627824,0.00082647306,0.00016514398,0.00027640883,0.0007508019,0.0028623894,0.08021998],"study_design_scores_gemma":[0.00006440345,0.00011389864,0.0024282092,0.94183385,0.018553108,0.00028108095,0.00069609703,0.000059458744,0.00022313058,0.0006249982,0.035095945,0.000025744075],"about_ca_topic_score_codex":0.007647828,"about_ca_topic_score_gemma":0.031964928,"teacher_disagreement_score":0.017252557,"about_ca_system_score_codex":0.0031730598,"about_ca_system_score_gemma":0.02076023,"threshold_uncertainty_score":0.08657181},"labels":[],"label_agreement":null},{"id":"W4411306282","doi":"10.1007/s41666-025-00204-w","title":"LongHealth: A Question Answering Benchmark with Long Clinical Documents","year":2025,"lang":"en","type":"article","venue":"Journal of Healthcare Informatics Research","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Technische Universität München","keywords":"Benchmark (surveying); Computer science; Sorting; Identification (biology); Unstructured data; Data science; Information retrieval; Data mining; Artificial intelligence; Natural language processing; Big data; Programming language","score_opus":0.10737496601511887,"score_gpt":0.5006962605306308,"score_spread":0.3933212945155119,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411306282","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45430735,0.025243765,0.13725403,0.01717604,0.0033369127,0.0031770114,0.20494635,0.11511905,0.0394395],"genre_scores_gemma":[0.5021556,0.0016775584,0.15809387,0.003874267,0.0005334038,0.0012485259,0.32147416,0.002347391,0.008595166],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99283606,0.003450789,0.0007891582,0.0014911232,0.0011516169,0.00028119417],"domain_scores_gemma":[0.9748908,0.018632982,0.0005593622,0.0020607254,0.003026236,0.0008298088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00692197,0.0021067094,0.0008129067,0.0024042695,0.0009345267,0.0022901683,0.0035809458,0.0034641086,0.010116554],"category_scores_gemma":[0.0388884,0.00048048844,0.001249588,0.0019890866,0.0009913255,0.002802691,0.0024474685,0.0020162582,0.0055491407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0047114617,0.002212934,0.020270284,0.006161406,0.0007657226,0.0018732949,0.0015812614,0.113842584,0.011851018,0.007789371,0.41111276,0.41782787],"study_design_scores_gemma":[0.0023235995,0.0027089573,0.02112528,0.0012277993,0.00033650527,0.0024917081,0.0023173809,0.6862083,0.038977534,0.028032923,0.2138954,0.00035467342],"about_ca_topic_score_codex":0.01651224,"about_ca_topic_score_gemma":0.01810037,"teacher_disagreement_score":0.01651224,"about_ca_system_score_codex":0.0023604915,"about_ca_system_score_gemma":0.0026426355,"threshold_uncertainty_score":0.036607325},"labels":[],"label_agreement":null},{"id":"W4411332331","doi":"10.18280/mmep.120517","title":"Learning-Based Information Extraction to Obtain Prominent Named Entities in Indonesian Court Decision Documents","year":2025,"lang":"en","type":"article","venue":"Mathematical Modelling and Engineering Problems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Indonesian; Computer science; Natural language processing; Information extraction; Information retrieval; Artificial intelligence; Linguistics; Philosophy","score_opus":0.009365084124116418,"score_gpt":0.2276536113526836,"score_spread":0.2182885272285672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411332331","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26798445,0.0035969468,0.69501126,0.0018779847,0.00028844026,0.00051499,0.009975053,0.007291508,0.013459434],"genre_scores_gemma":[0.6815966,0.0014423206,0.29561254,0.00017295075,0.00013840494,0.0001798129,0.015964106,0.00013168836,0.004761658],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99933356,0.00014594993,0.00009728066,0.00021249872,0.00013156275,0.000079109814],"domain_scores_gemma":[0.99842095,0.0009602135,0.0001776662,0.00015979324,0.00024603168,0.000035272795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008004027,0.00061503175,0.00047868266,0.003956228,0.0005763873,0.0010802678,0.00063960487,0.0007260418,0.001608226],"category_scores_gemma":[0.0035214082,0.00024274517,0.00074298284,0.0030267239,0.00034964003,0.0027863707,0.00084905024,0.0009902437,0.0017101022],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023252887,0.00020738486,0.011746364,0.0006035106,0.00009601179,0.0012037955,0.00077671086,0.030766243,0.019092541,0.01028591,0.017511956,0.90747714],"study_design_scores_gemma":[0.000035896708,0.000104025516,0.01652893,0.00018270475,0.00018364108,0.0009558613,0.0009704567,0.88230366,0.044963352,0.020974815,0.032720085,0.000076601485],"about_ca_topic_score_codex":0.0055315983,"about_ca_topic_score_gemma":0.009034629,"teacher_disagreement_score":0.0055315983,"about_ca_system_score_codex":0.0007304595,"about_ca_system_score_gemma":0.0016192171,"threshold_uncertainty_score":0.010998845},"labels":[],"label_agreement":null},{"id":"W4411346145","doi":"10.1145/3743676","title":"Transphobia Is in the Eye of the Prompter: Trans-Centered Perspectives on Large Language Models","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Computer-Human Interaction","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Transphobia; Psychology; Optometry; Ophthalmology; Artificial intelligence; Computer science; Medicine; Psychoanalysis","score_opus":0.024914344486874972,"score_gpt":0.3046337098118638,"score_spread":0.2797193653249888,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411346145","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40199727,0.001571616,0.4238826,0.06365617,0.00041953046,0.00033985986,0.00036308097,0.000950307,0.10681958],"genre_scores_gemma":[0.9727726,0.00023446932,0.020756084,0.0016910394,0.000081741,0.00013622285,0.00007933807,0.00037830914,0.0038703145],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9629857,0.033113047,0.00042818848,0.0012994374,0.0013632007,0.0008104539],"domain_scores_gemma":[0.9180103,0.06665064,0.003106933,0.006809157,0.0034769282,0.0019460156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031362504,0.0010516028,0.0006376241,0.0020059547,0.0053510773,0.014183023,0.0015257493,0.0029190413,0.006679848],"category_scores_gemma":[0.05426211,0.0007621501,0.00073392637,0.0011921017,0.01605532,0.027727803,0.010855773,0.0056253793,0.0012419063],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017435288,0.00007283117,0.006156572,0.00028239228,0.00002687611,0.0004473657,0.74652463,0.0007755164,0.004521699,0.21362907,0.0030287413,0.024359947],"study_design_scores_gemma":[0.000044186345,0.00017179422,0.003957929,0.000720757,0.00006463798,0.00084186636,0.6091724,0.018164057,0.004747022,0.22488134,0.13710506,0.00012894935],"about_ca_topic_score_codex":0.004074096,"about_ca_topic_score_gemma":0.005419415,"teacher_disagreement_score":0.031362504,"about_ca_system_score_codex":0.004613602,"about_ca_system_score_gemma":0.0029236167,"threshold_uncertainty_score":0.16586274},"labels":[],"label_agreement":null},{"id":"W4411389650","doi":"10.21203/rs.3.rs-6892199/v1","title":"Speaking in Feelings: Facilitating Human Emotion Communication through Analogy by Large Language Models","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Analogy; Feeling; Psychology; Human communication; Communication; Linguistics; Cognitive psychology; Cognitive science; Computer science; Social psychology; Philosophy","score_opus":0.10682082857609647,"score_gpt":0.43552638681337275,"score_spread":0.32870555823727626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411389650","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09316634,0.00048768084,0.8980461,0.0010702307,0.00012672185,0.000121187586,0.0001669264,0.0017374103,0.0050775516],"genre_scores_gemma":[0.8295333,0.0003033059,0.16616228,0.00024718398,0.00012315201,0.00019700498,0.00034922583,0.00029417776,0.0027904075],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99819833,0.0013006388,0.000042902622,0.00028072906,0.000118044845,0.000059315593],"domain_scores_gemma":[0.9937604,0.0053119054,0.00017074148,0.00045958895,0.00016904883,0.0001283502],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00219755,0.0005178242,0.0006512903,0.0004214387,0.0006298144,0.0016605834,0.0010483719,0.0010153456,0.0054212636],"category_scores_gemma":[0.014275035,0.00041987278,0.0008216814,0.0004796085,0.0006492953,0.0045181494,0.0020897505,0.0021568472,0.0012958723],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026051064,0.0015077271,0.0068844566,0.0012636855,0.00045151004,0.0006992104,0.012318234,0.14059412,0.099109635,0.12033405,0.015286988,0.5989452],"study_design_scores_gemma":[0.000104713545,0.00018187005,0.0012395089,0.00003229096,0.00013556163,0.00015394845,0.0005417027,0.90453327,0.010013535,0.0777828,0.005235413,0.00004542764],"about_ca_topic_score_codex":0.00072323467,"about_ca_topic_score_gemma":0.0007992721,"teacher_disagreement_score":0.0054212636,"about_ca_system_score_codex":0.0003387874,"about_ca_system_score_gemma":0.00044735335,"threshold_uncertainty_score":0.018135905},"labels":[],"label_agreement":null},{"id":"W4411391110","doi":"10.1177/17470218251353509","title":"Updating social knowledge via episodic memory prediction errors","year":2025,"lang":"en","type":"article","venue":"Quarterly Journal of Experimental Psychology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Episodic memory; Cognitive psychology; Semantic memory; Psychology; Computer science; Cognition; Neuroscience","score_opus":0.023464562515547738,"score_gpt":0.3506558394004344,"score_spread":0.32719127688488664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411391110","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9038739,0.0002944171,0.08718289,0.0005424042,0.000050772396,0.000095252835,0.00025595215,0.0004068761,0.0072975797],"genre_scores_gemma":[0.9928532,0.000099039986,0.006182929,0.0000382145,0.000016914359,0.000022450517,0.00016312156,0.000018763423,0.00060529145],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99859124,0.00041843794,0.00011693454,0.00041595878,0.00034842853,0.00010897772],"domain_scores_gemma":[0.97135025,0.016863909,0.005869068,0.0039196345,0.0014877836,0.0005093317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028127693,0.00055705936,0.0004993446,0.00090428727,0.0002960892,0.0024201216,0.0009101021,0.0007404277,0.002342184],"category_scores_gemma":[0.045193736,0.0006248587,0.0003486988,0.0005943414,0.0008330065,0.003937793,0.0014306586,0.001212496,0.00033695716],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023358809,0.0010310726,0.220218,0.00036328397,0.00067690975,0.0007747664,0.0041149002,0.090145856,0.035721853,0.024906999,0.0023826235,0.61732787],"study_design_scores_gemma":[0.00018782263,0.00092861004,0.20003143,0.0002129089,0.00039597973,0.0008308505,0.0011292499,0.64696133,0.028665746,0.11723097,0.003250775,0.00017425824],"about_ca_topic_score_codex":0.0027726989,"about_ca_topic_score_gemma":0.002204749,"teacher_disagreement_score":0.0028127693,"about_ca_system_score_codex":0.00064393936,"about_ca_system_score_gemma":0.00058753876,"threshold_uncertainty_score":0.014875531},"labels":[],"label_agreement":null},{"id":"W4411403346","doi":"10.1145/3744746","title":"A Comprehensive Overview of Large Language Models","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Intelligent Systems and Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":488,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Computer science; Context (archaeology); Benchmarking; Frontier; Data science; Engineering ethics; Management science; Political science; Engineering; Management","score_opus":0.041916858637249886,"score_gpt":0.307768274351576,"score_spread":0.26585141571432613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411403346","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027032692,0.17498478,0.7742202,0.008656428,0.0011088821,0.00032368448,0.00836263,0.005980137,0.023659931],"genre_scores_gemma":[0.116883256,0.27176562,0.5455529,0.0042337817,0.0060451096,0.0019747405,0.032051522,0.0025580348,0.018934932],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99738544,0.0010457587,0.00029512288,0.0004219268,0.0007491369,0.00010264032],"domain_scores_gemma":[0.9939644,0.0043456266,0.00027877107,0.0006233522,0.00067247957,0.00011542134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035038546,0.0017370526,0.0014591195,0.0044655725,0.0007505256,0.0044085854,0.002607852,0.0017404908,0.010180098],"category_scores_gemma":[0.012138018,0.00097230775,0.002192413,0.0054227468,0.0008275392,0.0065181716,0.0020880667,0.0029543764,0.007107844],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010972584,0.00011838186,0.0021259016,0.004219814,0.00039591047,0.0003827566,0.00048092206,0.051077504,0.002023827,0.1914044,0.090913765,0.65674704],"study_design_scores_gemma":[0.000023253862,0.000084636056,0.0014110126,0.0014233928,0.00018286383,0.0006307563,0.0001858462,0.19381233,0.001348266,0.32290113,0.47788337,0.000113182025],"about_ca_topic_score_codex":0.005581778,"about_ca_topic_score_gemma":0.0061090426,"teacher_disagreement_score":0.010180098,"about_ca_system_score_codex":0.0017571669,"about_ca_system_score_gemma":0.0034312129,"threshold_uncertainty_score":0.03405577},"labels":[],"label_agreement":null},{"id":"W4411447376","doi":"10.1109/icsc64641.2025.00032","title":"MediTriR: A Triple-Driven Approach to Retrieval-Augmented Generation for Medical Question Answering Tasks","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Question answering; Computer science; Information retrieval; Artificial intelligence; Natural language processing","score_opus":0.03046773196624122,"score_gpt":0.3030382831378181,"score_spread":0.2725705511715769,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411447376","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005981646,0.0006581379,0.9654477,0.0007388396,0.00017352105,0.00075451564,0.003160811,0.021126295,0.0019584773],"genre_scores_gemma":[0.093699865,0.0004160155,0.88975286,0.0006347788,0.00013014728,0.00085128745,0.011002278,0.001131636,0.0023810926],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9966029,0.0016361391,0.00025078925,0.0006372718,0.0007420966,0.00013085398],"domain_scores_gemma":[0.9947259,0.0032358356,0.0002037095,0.0009807755,0.00068149116,0.00017228839],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004351448,0.0015322989,0.0010866527,0.0030407882,0.00069873896,0.0015698711,0.0029690233,0.0017092826,0.006270327],"category_scores_gemma":[0.013112951,0.00062674016,0.002893916,0.0016004742,0.00078455155,0.0029147437,0.003662386,0.0023070055,0.003616022],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009111493,0.00096156634,0.0036123313,0.0018617898,0.0005555331,0.0011279372,0.0017915444,0.12313085,0.032211244,0.035712313,0.07983436,0.7182894],"study_design_scores_gemma":[0.00017416055,0.00022384997,0.00060928747,0.0001051971,0.00015369058,0.00042791365,0.0003255257,0.87132,0.016653817,0.07200309,0.037891425,0.00011207874],"about_ca_topic_score_codex":0.0061520212,"about_ca_topic_score_gemma":0.010737051,"teacher_disagreement_score":0.006270327,"about_ca_system_score_codex":0.0010301435,"about_ca_system_score_gemma":0.0021656095,"threshold_uncertainty_score":0.023012936},"labels":[],"label_agreement":null},{"id":"W4411447528","doi":"10.1109/icsc64641.2025.00011","title":"TriRAG: Enhancing Retrieval-Augmented Generation Method with Triple-Based Knowledge Graphs for Improved Question Answering","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Question answering; Computer science; Information retrieval; Knowledge graph; Artificial intelligence","score_opus":0.023017567291996256,"score_gpt":0.31437519392034546,"score_spread":0.2913576266283492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411447528","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018398775,0.00064543437,0.9565873,0.00044332969,0.00014366701,0.00035938722,0.0008978702,0.020635365,0.0018888449],"genre_scores_gemma":[0.23898578,0.00040164866,0.74696046,0.0007355106,0.00011435344,0.00060967944,0.0061281943,0.0011919807,0.0048723263],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998245,0.00073433266,0.00011419464,0.00044655363,0.00036128526,0.00009857637],"domain_scores_gemma":[0.9973062,0.0014420976,0.000107290616,0.00061769754,0.0004397875,0.0000870113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023014008,0.001246815,0.0009864513,0.0021874781,0.00056536216,0.0010180072,0.0018409841,0.0014078618,0.00455521],"category_scores_gemma":[0.0075191483,0.0003683701,0.0018824645,0.0014643827,0.00069654756,0.003257112,0.0019052541,0.0017415358,0.0027743555],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041157624,0.00053827505,0.0025615252,0.0008262773,0.00021067598,0.0005494111,0.0010541839,0.058832377,0.042067382,0.01705057,0.03789281,0.83800495],"study_design_scores_gemma":[0.00014628416,0.0002491937,0.00081464567,0.000043425665,0.00012666735,0.00041484804,0.00027156496,0.92836976,0.023813223,0.024980566,0.020682136,0.00008755826],"about_ca_topic_score_codex":0.0050418535,"about_ca_topic_score_gemma":0.007019293,"teacher_disagreement_score":0.0050418535,"about_ca_system_score_codex":0.0006322604,"about_ca_system_score_gemma":0.0013571001,"threshold_uncertainty_score":0.015238643},"labels":[],"label_agreement":null},{"id":"W4411449748","doi":"10.1145/3715735","title":"Hallucination Detection in Large Language Models with Metamorphic Relations","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on software engineering.","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Calgary","funders":"","keywords":"Recall; Computer science; Margin (machine learning); Cognitive psychology; Psychology; Artificial intelligence; Machine learning","score_opus":0.00800568433846788,"score_gpt":0.20890885471749607,"score_spread":0.2009031703790282,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411449748","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21072254,0.0029204183,0.7291935,0.0013826889,0.00016596388,0.00027100442,0.0020285307,0.050792318,0.002523038],"genre_scores_gemma":[0.8225754,0.00040258074,0.17017597,0.0009175448,0.000077481665,0.00014293403,0.003407765,0.000859108,0.0014413692],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9957956,0.0016460613,0.00034602688,0.0010344423,0.0010014945,0.0001762769],"domain_scores_gemma":[0.9839357,0.011386541,0.0012417856,0.0021663848,0.0010183908,0.00025121096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044403435,0.0016290031,0.00096973946,0.0016024695,0.0005243564,0.0019141111,0.0017742619,0.0015276129,0.0010413295],"category_scores_gemma":[0.027853578,0.00058618403,0.0014371332,0.0008889667,0.0010615648,0.0035927831,0.0030999796,0.0019480314,0.0009720192],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012957522,0.0004617297,0.038737785,0.0014161898,0.0010569214,0.002796693,0.0019496321,0.27758828,0.039259307,0.0066338177,0.021747962,0.60705596],"study_design_scores_gemma":[0.00004082602,0.000103991246,0.0016591571,0.000035742694,0.00006977429,0.00049041264,0.00022465237,0.9760089,0.00959558,0.009499172,0.0022344908,0.000037310532],"about_ca_topic_score_codex":0.0049680877,"about_ca_topic_score_gemma":0.007165468,"teacher_disagreement_score":0.0049680877,"about_ca_system_score_codex":0.0009038101,"about_ca_system_score_gemma":0.0010501721,"threshold_uncertainty_score":0.023483038},"labels":[],"label_agreement":null},{"id":"W4411531987","doi":"10.1007/978-3-031-96235-6_22","title":"Enhancing Answer Reliability Through Inter-Model Consensus of Large Language Models","year":2025,"lang":"en","type":"book-chapter","venue":"IFIP advances in information and communication technology","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Reliability (semiconductor); Computer science; Natural language processing; Physics","score_opus":0.011140105632741579,"score_gpt":0.2817262020971429,"score_spread":0.2705860964644013,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411531987","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011027926,0.00015150967,0.9848475,0.0004007727,0.000052499647,0.00005881033,0.000084494786,0.0021530776,0.0012234264],"genre_scores_gemma":[0.4803975,0.00029918534,0.51168627,0.00043645294,0.00017628534,0.00029716606,0.00085102656,0.0014559473,0.004400139],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99076146,0.0046086633,0.0004759776,0.0013666906,0.0022858095,0.00050130126],"domain_scores_gemma":[0.9504877,0.034456927,0.0015478742,0.008176319,0.004765736,0.00056548225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009023617,0.0011892924,0.0020803707,0.0013706151,0.001108074,0.002777049,0.00341814,0.002082516,0.0047988617],"category_scores_gemma":[0.05088921,0.0012300835,0.0017534448,0.0013306633,0.0013920373,0.007645885,0.0054840553,0.0045219166,0.002376008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010996393,0.0005202029,0.002140804,0.0006131365,0.00036446404,0.00031892027,0.001671386,0.40222666,0.032720204,0.12293448,0.016194506,0.41919547],"study_design_scores_gemma":[0.00003350288,0.000055056,0.00009160047,0.000013967025,0.0000505495,0.000031770178,0.000078646415,0.92555624,0.0056262747,0.06679585,0.0016495486,0.000016895348],"about_ca_topic_score_codex":0.0026153417,"about_ca_topic_score_gemma":0.0040252437,"teacher_disagreement_score":0.009023617,"about_ca_system_score_codex":0.0013408305,"about_ca_system_score_gemma":0.0022782127,"threshold_uncertainty_score":0.04772204},"labels":[],"label_agreement":null},{"id":"W4411630053","doi":"10.18653/v1/2024.eacl-long.175","title":"Gradient-Based Language Model Red Teaming","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Natural language processing; Artificial intelligence","score_opus":0.019019995887278443,"score_gpt":0.2566363238250168,"score_spread":0.23761632793773835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411630053","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017738415,0.00034277636,0.96781766,0.0005388707,0.00015834128,0.000118117845,0.0001686293,0.009456862,0.0036603212],"genre_scores_gemma":[0.65259415,0.00024012345,0.32906497,0.0013388873,0.00016624539,0.00040137535,0.0007939941,0.0022644761,0.013135727],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99813867,0.0008232837,0.000067445515,0.0005223566,0.00028311874,0.00016510456],"domain_scores_gemma":[0.9958341,0.0026340976,0.00023449582,0.00059677317,0.00048928463,0.00021118387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026056094,0.0017668717,0.0012339114,0.00063271704,0.0006307305,0.0011619234,0.0023657975,0.0018038445,0.0069168326],"category_scores_gemma":[0.009398554,0.0006925542,0.0011895926,0.0003804849,0.0014164567,0.0024388616,0.0024353177,0.0034998197,0.0034329484],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042171156,0.00019314712,0.0014337933,0.00030902756,0.00010867903,0.00031476998,0.0005520588,0.67916805,0.012822217,0.024209945,0.017827239,0.26263937],"study_design_scores_gemma":[0.000017457107,0.00003809363,0.000048633403,0.000008364118,0.000007872498,0.000028681752,0.000015812282,0.98841864,0.0019383527,0.0084371865,0.0010308262,0.000010094566],"about_ca_topic_score_codex":0.0033958135,"about_ca_topic_score_gemma":0.00577427,"teacher_disagreement_score":0.0069168326,"about_ca_system_score_codex":0.0011797975,"about_ca_system_score_gemma":0.0012654514,"threshold_uncertainty_score":0.02313906},"labels":[],"label_agreement":null},{"id":"W4411630208","doi":"10.18653/v1/2024.eacl-long.105","title":"Learning to Retrieve In-Context Examples for Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Context (archaeology); Natural language processing; Artificial intelligence; Context model; Information retrieval; History","score_opus":0.038251800548082554,"score_gpt":0.28693300571569935,"score_spread":0.2486812051676168,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411630208","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.109520115,0.0021628044,0.8666925,0.00083262153,0.00011973352,0.0002457672,0.0006579098,0.01723259,0.0025360584],"genre_scores_gemma":[0.60948265,0.0006922686,0.3820277,0.000569336,0.00013833109,0.00031021697,0.0022554449,0.0008764971,0.0036475658],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99892247,0.00039800015,0.00006797975,0.0003523877,0.00016493807,0.00009423681],"domain_scores_gemma":[0.9966456,0.001988491,0.00020639178,0.00074624404,0.00028737253,0.00012588834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023309384,0.001963935,0.0014793124,0.0011549444,0.0007318939,0.0011610483,0.0029869585,0.0021769342,0.00329177],"category_scores_gemma":[0.011033434,0.00094859785,0.0011369977,0.00088480656,0.0008292265,0.0040839794,0.0027161902,0.0035315375,0.0023514703],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055893866,0.00047230668,0.005617115,0.00049974996,0.0002430553,0.00029873627,0.00043225285,0.32063177,0.016354885,0.0073161884,0.01302882,0.6345462],"study_design_scores_gemma":[0.000031557483,0.000080017686,0.00029106584,0.000015826989,0.000030139912,0.00007332228,0.000038540453,0.98658025,0.005058788,0.006803541,0.0009810618,0.00001587559],"about_ca_topic_score_codex":0.005179428,"about_ca_topic_score_gemma":0.0143810585,"teacher_disagreement_score":0.005179428,"about_ca_system_score_codex":0.0010485911,"about_ca_system_score_gemma":0.0011174808,"threshold_uncertainty_score":0.012327373},"labels":[],"label_agreement":null},{"id":"W4411644201","doi":"10.1145/3742423","title":"Utilizing Large Language Model for Conversational Information Seeking via Dual-Query Generation and Joint-Encoding","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; Western University","funders":"","keywords":"Computer science; Encoding (memory); Joint (building); Dual (grammatical number); Natural language processing; Language model; Artificial intelligence; Linguistics; Engineering","score_opus":0.033160082311940495,"score_gpt":0.2612312934812345,"score_spread":0.22807121116929402,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411644201","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022536172,0.00061995455,0.96868616,0.00063985796,0.00007668478,0.00029902032,0.00052271,0.0044494644,0.0021698738],"genre_scores_gemma":[0.44851735,0.00046221874,0.54139316,0.00058362115,0.00013252693,0.00052114663,0.00253908,0.00061922765,0.0052316664],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969284,0.0013591577,0.00020208841,0.0005921909,0.00071879936,0.00019948235],"domain_scores_gemma":[0.9964994,0.001868005,0.00015849016,0.0007460819,0.0005900986,0.00013791285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026153713,0.0012055953,0.0010342081,0.0012400481,0.0007444914,0.0015071844,0.0019417503,0.001261536,0.0027899768],"category_scores_gemma":[0.009566088,0.0004763349,0.0015320698,0.0011541685,0.0009107739,0.0042846305,0.002432634,0.0019918724,0.0016801778],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010681177,0.0007171085,0.004687105,0.0009368248,0.00022846904,0.00095075765,0.0031207516,0.14106418,0.06954293,0.0662528,0.02683374,0.6845972],"study_design_scores_gemma":[0.000058723675,0.00013138117,0.00025567733,0.000015512574,0.00006891129,0.00025763692,0.00029938325,0.9576972,0.011722473,0.022756673,0.006676468,0.000059968188],"about_ca_topic_score_codex":0.0067344657,"about_ca_topic_score_gemma":0.008161385,"teacher_disagreement_score":0.0067344657,"about_ca_system_score_codex":0.0012357198,"about_ca_system_score_gemma":0.001920524,"threshold_uncertainty_score":0.013831615},"labels":[],"label_agreement":null},{"id":"W4411806137","doi":"10.2196/70706","title":"Extracting Knowledge From Scientific Texts on Patient-Derived Cancer Models Using Large Language Models: Algorithm Development and Validation Study","year":2025,"lang":"en","type":"article","venue":"JMIR Bioinformatics and Biotechnology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Cancer Institute","keywords":"Computer science; Natural language processing; Algorithm; Artificial intelligence","score_opus":0.03245524288182356,"score_gpt":0.2980792810151806,"score_spread":0.26562403813335705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411806137","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21128152,0.0030455142,0.7444595,0.0020675848,0.00022937042,0.000990857,0.0050622732,0.030159349,0.002703941],"genre_scores_gemma":[0.39936796,0.0008141652,0.5838881,0.00046339817,0.00010425937,0.0006965506,0.012722893,0.00040773264,0.0015349614],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968454,0.001576648,0.00029771257,0.00075782347,0.00037913508,0.00014332507],"domain_scores_gemma":[0.9665914,0.02890348,0.0007239589,0.0014325654,0.0020435483,0.00030501836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074308557,0.0015074328,0.00093373895,0.0024394745,0.0005702385,0.0018298066,0.0016893399,0.0019859693,0.0038149978],"category_scores_gemma":[0.031293713,0.000508993,0.0014478196,0.0017931892,0.00049981155,0.002185933,0.0017684132,0.0026051025,0.0019765848],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000894425,0.0007160304,0.016289396,0.00096104224,0.00044094026,0.00062892673,0.0004917585,0.30474457,0.00480359,0.0034284426,0.01668393,0.64991707],"study_design_scores_gemma":[0.00008632615,0.00008316986,0.0011472023,0.000047477824,0.00004900629,0.0001247065,0.00010397422,0.9921159,0.002067307,0.0022715225,0.0018899618,0.000013389131],"about_ca_topic_score_codex":0.0070325895,"about_ca_topic_score_gemma":0.009546185,"teacher_disagreement_score":0.0074308557,"about_ca_system_score_codex":0.0015843984,"about_ca_system_score_gemma":0.0029434243,"threshold_uncertainty_score":0.039298594},"labels":[],"label_agreement":null},{"id":"W4411884041","doi":"10.1038/s41467-025-60762-w","title":"Extracting circumstances of Covid-19 transmission from free text with large language models","year":2025,"lang":"en","type":"article","venue":"Nature Communications","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Jewish General Hospital","funders":"Institut Pasteur; Agence Nationale de la Recherche","keywords":"Coronavirus disease 2019 (COVID-19); Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); 2019-20 coronavirus outbreak; Transmission (telecommunications); Computer science; Pandemic; Virology; Computational biology; Medicine; Biology; Outbreak; Telecommunications; Disease; Infectious disease (medical specialty)","score_opus":0.027900409383293405,"score_gpt":0.3268215091203649,"score_spread":0.2989210997370715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411884041","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5443839,0.0033994443,0.36859944,0.005470646,0.0009216347,0.0013381524,0.061242163,0.008114026,0.0065305806],"genre_scores_gemma":[0.8129944,0.00067850616,0.13369384,0.0006713356,0.0004626388,0.0008705215,0.048484847,0.0002062158,0.0019376294],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99568045,0.0026360657,0.00043193455,0.00080828014,0.00028830508,0.00015497339],"domain_scores_gemma":[0.9758976,0.020362195,0.0015207424,0.00094796054,0.00095752534,0.00031394773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047742683,0.0015624558,0.0006042191,0.0037679868,0.0005199339,0.0022041814,0.0010012062,0.0014496945,0.0019421672],"category_scores_gemma":[0.019848678,0.0004633547,0.0016089734,0.0014863302,0.0004692676,0.003906059,0.0018890388,0.0018169977,0.0024318243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002235654,0.0012057286,0.2470668,0.004835603,0.0010764365,0.004097173,0.011122995,0.09143491,0.028436257,0.010364128,0.06419874,0.5339255],"study_design_scores_gemma":[0.00015549816,0.00047089974,0.05763054,0.00069070165,0.00035215722,0.0010541684,0.0056559024,0.85599446,0.008225872,0.029747484,0.039755136,0.00026716565],"about_ca_topic_score_codex":0.0037995598,"about_ca_topic_score_gemma":0.0060480204,"teacher_disagreement_score":0.0047742683,"about_ca_system_score_codex":0.00087195076,"about_ca_system_score_gemma":0.0010819151,"threshold_uncertainty_score":0.025249064},"labels":[],"label_agreement":null},{"id":"W4412014287","doi":"10.1038/s41598-025-09052-5","title":"Benchmarking pre-trained text embedding models in aligning built asset information","year":2025,"lang":"en","type":"article","venue":"Scientific Reports","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Benchmarking; Computer science; Embedding; Asset (computer security); Data science; Artificial intelligence; Information retrieval; Natural language processing; Business; Computer security","score_opus":0.01602443871516109,"score_gpt":0.27933546170641804,"score_spread":0.26331102299125697,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412014287","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47861266,0.0086661,0.4586688,0.0017069681,0.0018749672,0.0013228497,0.016818821,0.018815799,0.013512985],"genre_scores_gemma":[0.6072483,0.002257788,0.3008481,0.0005971245,0.00032190036,0.0009537547,0.07980311,0.0008790373,0.0070907827],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99742573,0.001215295,0.00025803782,0.00068772526,0.00027918856,0.00013398219],"domain_scores_gemma":[0.99376124,0.0039725257,0.00024462296,0.0009575198,0.00090630824,0.00015782662],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041036806,0.0019653467,0.0007677356,0.0030451948,0.00055680407,0.0017145813,0.0017231356,0.0014883371,0.002398471],"category_scores_gemma":[0.012745332,0.00031151087,0.001259347,0.0027660013,0.0006049572,0.0038384213,0.0016279923,0.0019346497,0.0034993177],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008245557,0.00116,0.012143417,0.0016794296,0.0005249424,0.00037974896,0.00073617167,0.22360978,0.011196134,0.0034747273,0.034819663,0.7094514],"study_design_scores_gemma":[0.00008294305,0.0003697489,0.0049647843,0.00019289932,0.0001437291,0.0002178585,0.0006102123,0.9593728,0.014262257,0.0050033205,0.014709307,0.00007014738],"about_ca_topic_score_codex":0.0059032487,"about_ca_topic_score_gemma":0.00831981,"teacher_disagreement_score":0.0059032487,"about_ca_system_score_codex":0.0009878832,"about_ca_system_score_gemma":0.0010905566,"threshold_uncertainty_score":0.021702588},"labels":[],"label_agreement":null},{"id":"W4412042432","doi":"10.1007/978-3-031-85747-8_1","title":"Introduction and Fundamentals","year":2025,"lang":"en","type":"book-chapter","venue":"Machine translation","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computer science; Computational linguistics; Artificial intelligence; Natural language processing","score_opus":0.01891034828964161,"score_gpt":0.2374222924443348,"score_spread":0.2185119441546932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412042432","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00079353707,0.023891311,0.10552647,0.0040684314,0.0069376696,0.00023012966,0.0018732648,0.002880897,0.8537983],"genre_scores_gemma":[0.00458223,0.012832885,0.027060231,0.0013039069,0.0022667788,0.00018206933,0.0020108505,0.0011342994,0.9486267],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995227,0.00006667214,0.000023626453,0.00010331693,0.00024576986,0.000037896876],"domain_scores_gemma":[0.99945384,0.00015905286,0.000017532318,0.00009658357,0.00022244931,0.000050522387],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00051398197,0.0012756998,0.00088072877,0.0025152436,0.0011956561,0.003495042,0.0013180025,0.0014643973,0.21894605],"category_scores_gemma":[0.0018455584,0.0005897212,0.00061307556,0.0031813856,0.0008053297,0.0041893832,0.0016301654,0.002331595,0.19117288],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025950567,0.00005361173,0.00007832052,0.00040190745,0.0000053513118,0.000068173875,0.00021248832,0.00059181,0.001416688,0.09825885,0.47948393,0.41940296],"study_design_scores_gemma":[0.0000018618485,0.000011277707,0.00006436402,0.000107945816,0.0000027724404,0.000108928034,0.000040835082,0.00035919354,0.0004080622,0.021685814,0.9772017,0.0000070700867],"about_ca_topic_score_codex":0.0013725151,"about_ca_topic_score_gemma":0.0022825794,"teacher_disagreement_score":0.21894605,"about_ca_system_score_codex":0.00081658555,"about_ca_system_score_gemma":0.0013146857,"threshold_uncertainty_score":0.7324475},"labels":[],"label_agreement":null},{"id":"W4412065378","doi":"10.1007/978-3-031-85747-8_11","title":"Remaining Issues for AI","year":2025,"lang":"en","type":"book-chapter","venue":"Machine translation","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Computational linguistics; Computer science; Artificial intelligence; Natural language processing","score_opus":0.03829745308203281,"score_gpt":0.29902514287934606,"score_spread":0.26072768979731326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412065378","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013638359,0.035256658,0.0627637,0.18294299,0.012052041,0.000057701236,0.00055229536,0.00078111363,0.7042297],"genre_scores_gemma":[0.08034073,0.04292456,0.05458593,0.03433999,0.023460658,0.0002934755,0.001631947,0.0011956396,0.761227],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.998108,0.0008257266,0.00007729641,0.0002895854,0.000529159,0.00017016186],"domain_scores_gemma":[0.99517906,0.0024916397,0.00009535757,0.0009042737,0.001054417,0.00027527392],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036625636,0.00089302385,0.0010358499,0.0017819979,0.0042203325,0.009619728,0.0023023775,0.0037001623,0.09431982],"category_scores_gemma":[0.009829224,0.000518015,0.00074827863,0.0028604115,0.005716541,0.020405885,0.002877181,0.007435495,0.031378914],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015803227,0.000033971326,0.00003879455,0.00012171449,0.0000065999525,0.000023303772,0.00020729522,0.00033709034,0.00009847785,0.7681242,0.16135342,0.069639385],"study_design_scores_gemma":[0.0000041737353,0.0000045652096,0.000041687254,0.00007799963,0.00000309678,0.0000375218,0.00019315685,0.0008451283,0.00007379738,0.7341616,0.26454982,0.000007364247],"about_ca_topic_score_codex":0.0072055473,"about_ca_topic_score_gemma":0.0054626777,"teacher_disagreement_score":0.09431982,"about_ca_system_score_codex":0.0039464133,"about_ca_system_score_gemma":0.0033308435,"threshold_uncertainty_score":0.31553125},"labels":[],"label_agreement":null},{"id":"W4412075512","doi":"10.1101/2025.07.06.25330972","title":"Benchmarking Multimodal Large Language Models for Forensic Science and Medicine: A Comprehensive Dataset and Evaluation Framework","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Loyalist College; Queen's University","funders":"","keywords":"Benchmarking; Computer science; Data science; Forensic science; Natural language processing; Artificial intelligence; Medicine; Business","score_opus":0.05501007735612667,"score_gpt":0.3614847540115834,"score_spread":0.30647467665545675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412075512","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63177663,0.014033191,0.13653474,0.0062374296,0.0010312367,0.0052801813,0.1671169,0.02185828,0.01613134],"genre_scores_gemma":[0.51093024,0.0015574178,0.120432734,0.0012562159,0.00021131386,0.0033308682,0.35816997,0.00093707885,0.0031741292],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9863528,0.008557905,0.0010558364,0.0019613681,0.0016869162,0.00038508463],"domain_scores_gemma":[0.9746595,0.01568617,0.0010587826,0.004510453,0.0032769463,0.00080801675],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020177158,0.003323047,0.0013051578,0.0057012285,0.0014028637,0.0030118406,0.005394757,0.0035781004,0.0036146664],"category_scores_gemma":[0.040213328,0.00060222845,0.00283007,0.0036777295,0.0019934904,0.003372334,0.0046125553,0.0031291381,0.002794935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024440002,0.0043386845,0.07568038,0.00424357,0.002474803,0.0013106947,0.0013526282,0.34665185,0.004636257,0.0073937816,0.20782766,0.34164563],"study_design_scores_gemma":[0.0007421466,0.0017654217,0.036276333,0.0012808904,0.00072561257,0.0012289847,0.0018186088,0.87413114,0.010472739,0.013772662,0.05744185,0.00034369095],"about_ca_topic_score_codex":0.022527877,"about_ca_topic_score_gemma":0.027134111,"teacher_disagreement_score":0.022527877,"about_ca_system_score_codex":0.0041400976,"about_ca_system_score_gemma":0.0032551885,"threshold_uncertainty_score":0.10670829},"labels":[],"label_agreement":null},{"id":"W4412163797","doi":"10.1158/1557-3265.aimachine-b017","title":"Abstract B017: Check: A hybrid continuous-learning framework for enhancing factual reliability in clinical language models","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reliability (semiconductor); Computer science; Medicine; Natural language processing; Psychology; Reliability engineering; Engineering","score_opus":0.18614910490314512,"score_gpt":0.5355942354535079,"score_spread":0.3494451305503628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163797","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03843705,0.0008612401,0.9496384,0.0011914858,0.0000738712,0.00021051895,0.0009434084,0.0074754055,0.0011686026],"genre_scores_gemma":[0.53114474,0.0003063349,0.46241426,0.00074146804,0.00014447745,0.0003158556,0.0027293635,0.00040439062,0.001799076],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99738616,0.001329813,0.0001613704,0.0005247611,0.00048895366,0.00010902584],"domain_scores_gemma":[0.98956746,0.0074115447,0.0005180098,0.0008805193,0.0012725194,0.00034992932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069326186,0.0009587612,0.00086169917,0.0017688771,0.00043111484,0.0018911499,0.002518308,0.0016285001,0.0030168865],"category_scores_gemma":[0.018948887,0.0004461548,0.00094277726,0.0008457469,0.00091184967,0.0019784311,0.0020949107,0.0022681896,0.00095119386],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076266844,0.0004930412,0.01056829,0.00036303268,0.00029820928,0.00024500568,0.00042367415,0.43949828,0.0068156547,0.009998263,0.012205098,0.5183288],"study_design_scores_gemma":[0.000018051362,0.000047425525,0.00024080972,0.000017833367,0.000013355148,0.000016252414,0.00001508531,0.9946561,0.00086214166,0.003341283,0.0007620268,0.0000096499125],"about_ca_topic_score_codex":0.008356071,"about_ca_topic_score_gemma":0.011135363,"teacher_disagreement_score":0.008356071,"about_ca_system_score_codex":0.0011424005,"about_ca_system_score_gemma":0.00235371,"threshold_uncertainty_score":0.03666359},"labels":[],"label_agreement":null},{"id":"W4412376911","doi":"10.1145/3726302.3730281","title":"Gosling Grows Up: Retrieval with Learned Dense and Sparse Representations Using Anserini","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Artificial intelligence; Information retrieval","score_opus":0.05878817003501889,"score_gpt":0.31238153571832655,"score_spread":0.25359336568330765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412376911","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06903275,0.0047270376,0.8066729,0.0019069465,0.0007235096,0.00044573584,0.0053103436,0.08862561,0.022555197],"genre_scores_gemma":[0.22888756,0.0016015139,0.7156649,0.0011008389,0.00035396652,0.00027497677,0.022973262,0.004280368,0.024862645],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998868,0.0002807805,0.00007942885,0.00027987504,0.00034870644,0.00014333591],"domain_scores_gemma":[0.9983753,0.00062805635,0.000058415924,0.00056881737,0.00025642998,0.00011290761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024224988,0.0013510353,0.0012888059,0.002277141,0.00085617637,0.0025985849,0.0018975268,0.0014581583,0.011196983],"category_scores_gemma":[0.008319035,0.00054133654,0.001232435,0.0022645881,0.0006903533,0.0055750427,0.0033379649,0.0018641443,0.0095368745],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011520024,0.00030430374,0.0018019404,0.0006717898,0.00021284241,0.00044813482,0.00071850343,0.037811685,0.030093903,0.02485695,0.11626525,0.7856627],"study_design_scores_gemma":[0.00017639958,0.00030336843,0.001122859,0.00008707331,0.00011018609,0.00041195002,0.00041074882,0.8663169,0.031177659,0.052309852,0.047456156,0.00011685403],"about_ca_topic_score_codex":0.009704512,"about_ca_topic_score_gemma":0.018677393,"teacher_disagreement_score":0.011196983,"about_ca_system_score_codex":0.0011204226,"about_ca_system_score_gemma":0.001316897,"threshold_uncertainty_score":0.037457585},"labels":[],"label_agreement":null},{"id":"W4412377058","doi":"10.1145/3726302.3730090","title":"The Great Nugget Recall: Automating Fact Extraction and RAG Evaluation with Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; National Institute of Standards and Technology; Ministry of Science and ICT, South Korea; Institute for Information and Communications Technology Promotion; Universitas Brawijaya","keywords":"Computer science; Recall; Natural language processing; Language model; Extraction (chemistry); Artificial intelligence; Programming language; Linguistics","score_opus":0.0236931126214044,"score_gpt":0.31025027733062893,"score_spread":0.28655716470922454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412377058","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.116586305,0.0015847246,0.8235809,0.00128641,0.00020565487,0.000968109,0.0024082977,0.04770422,0.0056754337],"genre_scores_gemma":[0.4496185,0.00022915035,0.53863853,0.0004822854,0.00008335891,0.0006310863,0.0057983557,0.0027025202,0.0018162187],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.95634395,0.028759636,0.002416485,0.0050493074,0.006550512,0.0008801983],"domain_scores_gemma":[0.90644187,0.06230911,0.0047297957,0.017123975,0.008164284,0.0012309473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.039057545,0.0021625904,0.0014320203,0.004772114,0.0014300423,0.0057215896,0.0032748918,0.0024647191,0.0029780779],"category_scores_gemma":[0.09082968,0.0013124563,0.0019269686,0.002103268,0.0017537319,0.009182416,0.005418842,0.003613322,0.001971903],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001223578,0.0008490197,0.034026243,0.0017334607,0.001205343,0.00062645966,0.006858351,0.084297314,0.046072144,0.0202893,0.033183683,0.76963514],"study_design_scores_gemma":[0.00022703606,0.00070509966,0.014473232,0.00029764566,0.00037966602,0.00050201826,0.0011375505,0.86099917,0.06387925,0.02427732,0.032799937,0.0003220626],"about_ca_topic_score_codex":0.0107606035,"about_ca_topic_score_gemma":0.018443134,"teacher_disagreement_score":0.039057545,"about_ca_system_score_codex":0.0026699114,"about_ca_system_score_gemma":0.003591433,"threshold_uncertainty_score":0.20655847},"labels":[],"label_agreement":null},{"id":"W4412377345","doi":"10.1145/3726302.3730159","title":"A Human-AI Comparative Analysis of Prompt Sensitivity in LLM-Based Relevance Judgment","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Relevance (law); Sensitivity (control systems); Computer science; Artificial intelligence; Political science; Engineering; Law","score_opus":0.0379838156217058,"score_gpt":0.3335992028802042,"score_spread":0.29561538725849845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412377345","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8912592,0.003391689,0.07879243,0.00058785523,0.0004922155,0.0010007655,0.001772526,0.0041016983,0.018601604],"genre_scores_gemma":[0.9550153,0.00031130176,0.039061386,0.00040521484,0.00012143907,0.0006512047,0.001991463,0.00055590295,0.0018866761],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9451596,0.038032524,0.0030739664,0.00688635,0.0060114674,0.0008361734],"domain_scores_gemma":[0.50976795,0.42824966,0.012479831,0.021788938,0.024443218,0.0032703893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.036863998,0.0010050922,0.0009425845,0.002866886,0.0010761997,0.002948158,0.0011565208,0.0016956966,0.0032494396],"category_scores_gemma":[0.31235388,0.0005604396,0.0006041109,0.001747735,0.001494123,0.0026874267,0.0036421798,0.0019788097,0.0020341151],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.018331258,0.0022578945,0.15570626,0.005745331,0.0010763617,0.0012300645,0.046120364,0.0274177,0.12090579,0.0067562065,0.032614704,0.581838],"study_design_scores_gemma":[0.0019885693,0.008620767,0.54481065,0.0015009649,0.000908679,0.0038563118,0.016297374,0.21031164,0.10523477,0.039820172,0.06498644,0.0016636475],"about_ca_topic_score_codex":0.0019027593,"about_ca_topic_score_gemma":0.0020075429,"teacher_disagreement_score":0.036863998,"about_ca_system_score_codex":0.0013932787,"about_ca_system_score_gemma":0.0010998092,"threshold_uncertainty_score":0.1949578},"labels":[],"label_agreement":null},{"id":"W4412378237","doi":"10.1145/3726302.3730244","title":"Response Quality Assessment for Retrieval-Augmented Generation via Conditional Conformal Factuality","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Universitas Brawijaya","keywords":"Computer science; Conformal map; Quality (philosophy); Information retrieval; Artificial intelligence; Mathematics; Physics","score_opus":0.08704555234092608,"score_gpt":0.39048012836480894,"score_spread":0.30343457602388285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412378237","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.120495565,0.0014488905,0.8429873,0.001304364,0.00029499113,0.000977094,0.0020322928,0.023544554,0.0069149956],"genre_scores_gemma":[0.7729212,0.00022948533,0.21811111,0.00051038485,0.00019155738,0.0005161298,0.0037857173,0.0010145352,0.0027198286],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98398733,0.008102507,0.0008531554,0.002589744,0.003919589,0.0005477223],"domain_scores_gemma":[0.9348242,0.042751532,0.0028068842,0.012426239,0.0062152953,0.0009758649],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014571449,0.0014508698,0.0014077842,0.0023000226,0.0009100031,0.0023128155,0.0021272798,0.0020361799,0.006667965],"category_scores_gemma":[0.08909014,0.0004564423,0.0011827023,0.0011326121,0.0014371325,0.0037298647,0.00384706,0.0020468067,0.0024177996],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002984524,0.0006897045,0.024668898,0.0015661481,0.00033829233,0.00059312675,0.002573661,0.11745042,0.035030868,0.030806797,0.027071571,0.7562259],"study_design_scores_gemma":[0.00018799325,0.00055844063,0.0052342387,0.000100139136,0.00014606943,0.00045035695,0.0003530046,0.92781365,0.02044015,0.03486518,0.009743341,0.0001074252],"about_ca_topic_score_codex":0.0025624083,"about_ca_topic_score_gemma":0.0027364306,"teacher_disagreement_score":0.014571449,"about_ca_system_score_codex":0.0013606449,"about_ca_system_score_gemma":0.0019127385,"threshold_uncertainty_score":0.07706207},"labels":[],"label_agreement":null},{"id":"W4412396369","doi":"10.1145/3726302.3730305","title":"Benchmarking LLM-based Relevance Judgment Methods","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Benchmarking; Relevance (law); Computer science; Data science; Artificial intelligence; Political science; Business","score_opus":0.02706635386643755,"score_gpt":0.3433562945497067,"score_spread":0.3162899406832691,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412396369","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34896043,0.017074589,0.52060604,0.0018264664,0.001977628,0.003987523,0.00946657,0.06485207,0.031248685],"genre_scores_gemma":[0.56853634,0.0010775313,0.40460247,0.0006342683,0.00035980673,0.0017196533,0.01574033,0.0015806861,0.0057488536],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97334087,0.01367105,0.0026804435,0.0035730624,0.005885104,0.00084946415],"domain_scores_gemma":[0.9480521,0.03298598,0.0017660036,0.0065500764,0.009298692,0.0013471553],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.022352947,0.0024343182,0.0013404475,0.0056792735,0.0013915509,0.0027521204,0.0030644073,0.003048455,0.0052064457],"category_scores_gemma":[0.08576729,0.00058070547,0.0017176404,0.0032305338,0.0010870793,0.0034479846,0.0041765543,0.0025991227,0.004009596],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034004184,0.0019602058,0.014805928,0.003646118,0.0009163238,0.0003372414,0.0010399764,0.124478884,0.015538951,0.006246357,0.054512624,0.773117],"study_design_scores_gemma":[0.0007109735,0.0012274203,0.007523361,0.00025924353,0.00018411945,0.00029091493,0.00045839394,0.93834424,0.023616457,0.0108181955,0.016391877,0.00017473624],"about_ca_topic_score_codex":0.0067426325,"about_ca_topic_score_gemma":0.00932659,"teacher_disagreement_score":0.97764707,"about_ca_system_score_codex":0.0025126894,"about_ca_system_score_gemma":0.0029137074,"threshold_uncertainty_score":0.118215084},"labels":[],"label_agreement":null},{"id":"W4412438306","doi":"10.1145/3730436.3730529","title":"Comparative Analysis of Large Language Models for Context-Aware Code Completion using SAFIM Framework","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Context (archaeology); Code (set theory); Programming language; Natural language processing","score_opus":0.07661870692166582,"score_gpt":0.3764386418616509,"score_spread":0.29981993493998504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412438306","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6425819,0.0055824374,0.25857642,0.0017185272,0.00045879255,0.0007222912,0.02408163,0.057948686,0.008329381],"genre_scores_gemma":[0.7291416,0.0012550568,0.18317561,0.00030915998,0.000115825926,0.0007549335,0.079413325,0.0027662686,0.003068084],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969258,0.0014839274,0.00020227837,0.0006437335,0.00054881285,0.00019542807],"domain_scores_gemma":[0.9885279,0.0079903845,0.00040001998,0.0015107492,0.0012335138,0.00033740568],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004699735,0.0015482148,0.00091843726,0.0026734828,0.0006654102,0.0015977778,0.0020949645,0.001103683,0.0019905602],"category_scores_gemma":[0.020669749,0.0004121609,0.0021569263,0.0015622823,0.00051454717,0.0024678926,0.0015692448,0.0015052231,0.001555145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025789875,0.0013372322,0.03823053,0.0026894498,0.0011402115,0.0006583681,0.0017979544,0.5279652,0.007403161,0.012567318,0.054959364,0.3486722],"study_design_scores_gemma":[0.000068030284,0.00022549137,0.0035191409,0.000083677,0.00010367277,0.00011280984,0.00034259295,0.9794709,0.0032064943,0.004327078,0.008499643,0.00004057171],"about_ca_topic_score_codex":0.019787638,"about_ca_topic_score_gemma":0.027245535,"teacher_disagreement_score":0.019787638,"about_ca_system_score_codex":0.001622828,"about_ca_system_score_gemma":0.002340188,"threshold_uncertainty_score":0.039344907},"labels":[],"label_agreement":null},{"id":"W4412445032","doi":"10.1109/access.2025.3589319","title":"Large Language Models in Transportation: A Comprehensive Bibliometric Analysis of Emerging Trends, Challenges, and Future Research","year":2025,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Data science; Bibliometrics; Regional science; Management science; Library science; Geography; Engineering","score_opus":0.11106527706161809,"score_gpt":0.4089443022322629,"score_spread":0.2978790251706448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412445032","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40348426,0.11660534,0.3171,0.044122588,0.00084137404,0.0010022322,0.03421936,0.003339479,0.079285294],"genre_scores_gemma":[0.855982,0.03187076,0.09334623,0.00073349423,0.001129737,0.00095644384,0.013108224,0.0003779821,0.002495165],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97913927,0.010363402,0.0015115049,0.0014233318,0.006989875,0.0005724969],"domain_scores_gemma":[0.9035795,0.07323042,0.0070493612,0.0064543793,0.008434559,0.0012518037],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.019898659,0.001069649,0.002058775,0.0628622,0.0024271624,0.013032426,0.0016681388,0.00129937,0.0033733677],"category_scores_gemma":[0.100395076,0.0005770502,0.0024920378,0.12432614,0.0018785791,0.017989356,0.0044772383,0.0014018913,0.0007165816],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016290162,0.00031156468,0.17529234,0.0059874924,0.0022124355,0.00039516395,0.006338202,0.038255632,0.0009781311,0.1363302,0.038093913,0.595642],"study_design_scores_gemma":[0.000077795434,0.00033021826,0.1730868,0.0042437417,0.0017158943,0.0012860748,0.021811577,0.24226098,0.0023111482,0.3615579,0.19087888,0.00043903335],"about_ca_topic_score_codex":0.01018414,"about_ca_topic_score_gemma":0.011110628,"teacher_disagreement_score":0.98010135,"about_ca_system_score_codex":0.004620538,"about_ca_system_score_gemma":0.00651975,"threshold_uncertainty_score":0.1052354},"labels":[],"label_agreement":null},{"id":"W4412446013","doi":"10.1109/icici65870.2025.11069785","title":"AI-Powered Resume Screening System Using NLP and Machine Learning","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Horizon College and Seminary","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Machine learning; Speech recognition","score_opus":0.03346483338106629,"score_gpt":0.26753765930031237,"score_spread":0.23407282591924608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412446013","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064968675,0.0007915206,0.562929,0.0013854443,0.0007188865,0.0020652579,0.02348295,0.31335634,0.030301908],"genre_scores_gemma":[0.35003483,0.0008274756,0.5136912,0.0011232464,0.0007279668,0.0031150463,0.0571329,0.0035036271,0.06984377],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99920064,0.00012716893,0.00008936468,0.00026760553,0.00026005527,0.000055193246],"domain_scores_gemma":[0.997355,0.0011153616,0.00023703864,0.00042455483,0.00064469955,0.00022343022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013198094,0.00082151574,0.0006533696,0.0027008257,0.00072183786,0.0013995044,0.001242976,0.0006187909,0.017212333],"category_scores_gemma":[0.0037783417,0.00027753226,0.0005209893,0.0012340089,0.00021706041,0.0017826249,0.0012317495,0.00077951537,0.018030442],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007626474,0.00067916256,0.011291258,0.001082726,0.000101600046,0.0007644941,0.0009822554,0.004149534,0.043874327,0.0029044764,0.15648721,0.7769204],"study_design_scores_gemma":[0.0003460329,0.00077065267,0.045492053,0.00030514938,0.00017923262,0.0010813667,0.0017314255,0.46625033,0.09421227,0.013288536,0.37604824,0.00029478263],"about_ca_topic_score_codex":0.0031716905,"about_ca_topic_score_gemma":0.0035077631,"teacher_disagreement_score":0.017212333,"about_ca_system_score_codex":0.00054533174,"about_ca_system_score_gemma":0.0011619631,"threshold_uncertainty_score":0.057581007},"labels":[],"label_agreement":null},{"id":"W4412496569","doi":"10.1007/978-3-031-98462-4_52","title":"From Recall to Reasoning: Automated Question Generation for Deeper Math Learning Through Large Language Models","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Recall; Automated reasoning; Artificial intelligence; Natural language processing; Cognitive science; Programming language; Mathematics education; Cognitive psychology; Mathematics; Psychology","score_opus":0.022786288251855216,"score_gpt":0.2863770605732907,"score_spread":0.2635907723214355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412496569","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020135159,0.00039071494,0.9453595,0.0010412261,0.00015410293,0.00024739202,0.0017196451,0.026037512,0.004914817],"genre_scores_gemma":[0.2700806,0.00028907627,0.71047896,0.00049576315,0.00015133437,0.0002747301,0.007877134,0.0018347944,0.008517652],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981647,0.000656132,0.000121364465,0.00056411733,0.00036782055,0.00012583974],"domain_scores_gemma":[0.9929529,0.0051564164,0.00018321077,0.0009193656,0.00061786006,0.00017018267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019400774,0.0014926048,0.0010919154,0.0015836986,0.0006796796,0.0028135318,0.002966801,0.0016747477,0.02101435],"category_scores_gemma":[0.010602954,0.001062021,0.002334566,0.00090897724,0.0009642573,0.0070047905,0.0038207793,0.0037587245,0.008161385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004942051,0.00042855428,0.0018405059,0.0004330286,0.00011802598,0.00024928522,0.00074130803,0.026496558,0.015683828,0.033864606,0.039843507,0.8798066],"study_design_scores_gemma":[0.00008948696,0.00009725193,0.0004266819,0.00007081176,0.00008837561,0.00012824922,0.00027713744,0.84106135,0.019648822,0.12489049,0.013183684,0.00003769552],"about_ca_topic_score_codex":0.0024070344,"about_ca_topic_score_gemma":0.0039029168,"teacher_disagreement_score":0.02101435,"about_ca_system_score_codex":0.0011604871,"about_ca_system_score_gemma":0.0014811694,"threshold_uncertainty_score":0.07030004},"labels":[],"label_agreement":null},{"id":"W4412496603","doi":"10.1007/978-3-031-98462-4_17","title":"An Emergent Bottom-Up Categorization of Students’ LLMs Usage in an Undergraduate Research Course","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Categorization; Course (navigation); Computer science; Mathematics education; Data science; Artificial intelligence; Psychology; Engineering","score_opus":0.05014343166486811,"score_gpt":0.36362174489695853,"score_spread":0.3134783132320904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412496603","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9879159,0.00016491907,0.0020757122,0.00041377818,0.00002635257,0.00001873527,0.00043082805,0.00013056639,0.00882308],"genre_scores_gemma":[0.9943025,0.00007507908,0.0015114049,0.00011326098,0.000018389655,0.000019333125,0.00061567273,0.000069531554,0.003274911],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.9987923,0.00030037222,0.00008114904,0.0002022937,0.00043092386,0.00019291203],"domain_scores_gemma":[0.992849,0.0037518477,0.0008754038,0.00036814145,0.0012137017,0.0009418369],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012888919,0.00017431233,0.00025178192,0.0023342196,0.00080861466,0.00302403,0.0005378095,0.00055320584,0.003862391],"category_scores_gemma":[0.009033733,0.00016291774,0.0002484453,0.0018752125,0.00046408185,0.0019647158,0.002150597,0.00077448564,0.0015510899],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005875066,0.0003142379,0.62545824,0.00023547547,0.000060728296,0.00047112274,0.094054826,0.00035700065,0.032561038,0.0064465962,0.010010451,0.22944275],"study_design_scores_gemma":[0.000009464886,0.00021454149,0.9134909,0.000120490884,0.0000252605,0.00040370843,0.05578402,0.0033482683,0.0029711607,0.003721498,0.019853406,0.00005736335],"about_ca_topic_score_codex":0.0023651978,"about_ca_topic_score_gemma":0.004915534,"teacher_disagreement_score":0.003862391,"about_ca_system_score_codex":0.000604818,"about_ca_system_score_gemma":0.0005502077,"threshold_uncertainty_score":0.012920976},"labels":[],"label_agreement":null},{"id":"W4412510611","doi":"10.1007/978-981-96-7008-6_31","title":"Role-Playing Based on Large Language Models via Style Extraction","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Style (visual arts); Extraction (chemistry); Computer science; Natural language processing; Artificial intelligence; Information retrieval; Art; Chromatography; Literature; Chemistry","score_opus":0.031609915465871566,"score_gpt":0.2982491380845073,"score_spread":0.26663922261863576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412510611","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010103415,0.0006812814,0.9712007,0.00026647907,0.00018342999,0.00017705077,0.00211307,0.010986211,0.0042883935],"genre_scores_gemma":[0.26263317,0.0010366291,0.71156955,0.00027298127,0.0002507849,0.00032321108,0.009955459,0.0024490661,0.01150913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987639,0.00039689543,0.00009403138,0.00035029056,0.00028472187,0.00011023471],"domain_scores_gemma":[0.9973832,0.0016348199,0.00010221342,0.00042135443,0.00035373858,0.00010469613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013215761,0.0016743491,0.0012097481,0.0021022263,0.000797439,0.0024975284,0.0017968401,0.0011186758,0.010468083],"category_scores_gemma":[0.004484006,0.00090092083,0.0022244398,0.0019676932,0.0004616833,0.0040069367,0.0014117864,0.0024756736,0.010814484],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006949421,0.00033714424,0.0022846854,0.0005460776,0.00019935638,0.00045743276,0.0005174887,0.030868486,0.043889254,0.02144926,0.02997912,0.8687768],"study_design_scores_gemma":[0.0000713031,0.00009262087,0.0008415347,0.00005238367,0.00012790265,0.0003392588,0.00017436878,0.92470497,0.018232618,0.03897281,0.016330356,0.00005978798],"about_ca_topic_score_codex":0.0030655686,"about_ca_topic_score_gemma":0.005895167,"teacher_disagreement_score":0.010468083,"about_ca_system_score_codex":0.0007077428,"about_ca_system_score_gemma":0.0011361121,"threshold_uncertainty_score":0.03501928},"labels":[],"label_agreement":null},{"id":"W4412544568","doi":"10.2139/ssrn.5163979","title":"Retrieval-Augmented Generation for Natural Language Processing: A Survey","year":2025,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Natural (archaeology); Computer science; Natural language processing; Information retrieval; Geography","score_opus":0.02733157135005905,"score_gpt":0.3038457472400465,"score_spread":0.27651417588998745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412544568","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012345049,0.11487778,0.8507466,0.0011796062,0.0004975816,0.0003185725,0.0012455749,0.009900916,0.008888321],"genre_scores_gemma":[0.18404606,0.10413455,0.68462163,0.0011961387,0.0015943064,0.0006559473,0.007921881,0.0024142994,0.013415192],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975973,0.0008196788,0.00020416724,0.00065000524,0.0006134098,0.00011545629],"domain_scores_gemma":[0.99447215,0.0035540191,0.00013072687,0.0011256742,0.0006359927,0.00008142818],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027460298,0.0011918071,0.0023417585,0.002797392,0.00058181986,0.0023970695,0.0037147736,0.0014748422,0.0077660717],"category_scores_gemma":[0.008488678,0.00084912713,0.0016477808,0.0039353734,0.000798591,0.004485541,0.0020453122,0.0013880434,0.0044944803],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013736656,0.00019775123,0.00061139395,0.0016493964,0.00007329248,0.00005650354,0.00015385408,0.009584587,0.003193516,0.010968436,0.012782338,0.9605915],"study_design_scores_gemma":[0.00014750891,0.00076150516,0.0037495883,0.0008358443,0.0004533812,0.0016296783,0.0005623301,0.5895298,0.030496484,0.12724195,0.24439955,0.00019241762],"about_ca_topic_score_codex":0.00379632,"about_ca_topic_score_gemma":0.00362838,"teacher_disagreement_score":0.0077660717,"about_ca_system_score_codex":0.0007666456,"about_ca_system_score_gemma":0.0017740806,"threshold_uncertainty_score":0.025980115},"labels":[],"label_agreement":null},{"id":"W4412563940","doi":"10.1007/s40593-025-00503-8","title":"Empowering Authorship with AI: a Novel Academic Writing Technology for Authorial Voice","year":2025,"lang":"en","type":"article","venue":"International Journal of Artificial Intelligence in Education","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Northern British Columbia","funders":"University of Otago","keywords":"Educational technology; Computer science; Multimedia; World Wide Web; Psychology; Mathematics education","score_opus":0.047176369009601336,"score_gpt":0.4123772131014436,"score_spread":0.3652008440918423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412563940","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023857279,0.0005592269,0.9127485,0.0045841956,0.0010079101,0.00023284339,0.000158262,0.0062301694,0.05062149],"genre_scores_gemma":[0.5406645,0.0005419759,0.415802,0.0011359053,0.0009108529,0.00040422723,0.00032125192,0.0009242905,0.039295048],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99353564,0.0030164686,0.0004392476,0.0009570104,0.0017455633,0.0003061081],"domain_scores_gemma":[0.97039557,0.018702354,0.0014423985,0.0056328997,0.0023492372,0.0014774583],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004509213,0.000657041,0.00063384447,0.0013430582,0.001966304,0.007609289,0.0024369494,0.0019307624,0.009865895],"category_scores_gemma":[0.024649367,0.00042487224,0.00084805815,0.001400403,0.0027936578,0.009606675,0.005987693,0.0028610427,0.0040598693],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003197855,0.0005197997,0.0025516066,0.00066212134,0.00008523568,0.00045333436,0.012621308,0.0054849847,0.0312051,0.4462645,0.022827312,0.4770049],"study_design_scores_gemma":[0.00012807173,0.00026080376,0.00085872784,0.00017183257,0.00014354903,0.0008451045,0.0031939386,0.15318064,0.035610765,0.6107698,0.1947146,0.0001221675],"about_ca_topic_score_codex":0.00031627252,"about_ca_topic_score_gemma":0.00045239693,"teacher_disagreement_score":0.009865895,"about_ca_system_score_codex":0.0006204563,"about_ca_system_score_gemma":0.001250297,"threshold_uncertainty_score":0.03300476},"labels":[],"label_agreement":null},{"id":"W4412620250","doi":"10.1007/978-981-96-7423-7_7","title":"Guided Response Generation for Conversational Recommender Systems","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"","keywords":"Computer science; Recommender system; Human–computer interaction; Artificial intelligence; Multimedia; Speech recognition; Information retrieval","score_opus":0.05895280063748062,"score_gpt":0.28790658831688254,"score_spread":0.22895378767940192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412620250","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010200463,0.0008913537,0.98069084,0.00038499228,0.0001946423,0.00021866863,0.00024525114,0.003360362,0.0038133815],"genre_scores_gemma":[0.36513928,0.00065349205,0.60824484,0.00046305996,0.00040983365,0.00071149477,0.0012856129,0.00068171293,0.022410695],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996716,0.002024283,0.0001284347,0.00043952596,0.00045447593,0.00023729399],"domain_scores_gemma":[0.9907873,0.0070875622,0.00016656454,0.0010042561,0.0007317547,0.00022256831],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037220677,0.0010754941,0.0018538198,0.0008683524,0.0012946867,0.0017635913,0.0029489934,0.002333076,0.012185634],"category_scores_gemma":[0.0127797695,0.0007530331,0.0009776414,0.0009752617,0.0006145712,0.0022394538,0.0023233404,0.002610725,0.0055494197],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013586507,0.0007434213,0.0011326419,0.00058874674,0.00021672233,0.00023833208,0.0010799575,0.13978274,0.011951431,0.03430566,0.037942197,0.77065957],"study_design_scores_gemma":[0.000058611324,0.00011103154,0.00014818847,0.000028419316,0.000038220714,0.000084857034,0.00008945268,0.96307427,0.002819278,0.028092096,0.005422345,0.000033187196],"about_ca_topic_score_codex":0.003484959,"about_ca_topic_score_gemma":0.0056049153,"teacher_disagreement_score":0.012185634,"about_ca_system_score_codex":0.0007366268,"about_ca_system_score_gemma":0.00091823906,"threshold_uncertainty_score":0.040764987},"labels":[],"label_agreement":null},{"id":"W4412673545","doi":"10.1145/3731120.3744605","title":"A Large-Scale Study of Relevance Assessments with Large Language Models Using UMBRELA","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Relevance (law); Computer science; Scale (ratio); Natural language processing; Artificial intelligence; Geography; Cartography","score_opus":0.02565721635038377,"score_gpt":0.3262908059187859,"score_spread":0.30063358956840214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412673545","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7337842,0.009759363,0.23274739,0.0012365425,0.0003752346,0.0012372832,0.0020558415,0.0119012445,0.0069028945],"genre_scores_gemma":[0.83511764,0.0004973364,0.15795964,0.00027803978,0.00019791648,0.00056189875,0.0028723418,0.000564832,0.0019503817],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95439935,0.036059566,0.0015224321,0.0034073666,0.0040497375,0.0005615861],"domain_scores_gemma":[0.80268997,0.16646604,0.004756533,0.015062835,0.009491268,0.0015332261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0382341,0.0022135687,0.0016963362,0.0051163086,0.0017179137,0.003127003,0.002410957,0.0017634402,0.0013879879],"category_scores_gemma":[0.13323542,0.0009335557,0.0013653712,0.0032453972,0.0012094076,0.0061556315,0.0023777583,0.003188063,0.0013322638],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035004949,0.0030551257,0.049775794,0.0030997957,0.0023281083,0.00047395742,0.0045290636,0.14916015,0.035912536,0.0051565254,0.026066734,0.7169418],"study_design_scores_gemma":[0.0004885642,0.0032499535,0.030422391,0.00018624464,0.0004244541,0.0004433342,0.0011126102,0.9179212,0.027799293,0.0060030553,0.011638987,0.00030994316],"about_ca_topic_score_codex":0.012141415,"about_ca_topic_score_gemma":0.021032594,"teacher_disagreement_score":0.0382341,"about_ca_system_score_codex":0.0021345569,"about_ca_system_score_gemma":0.0018978374,"threshold_uncertainty_score":0.20220369},"labels":[],"label_agreement":null},{"id":"W4412684160","doi":"10.1007/978-3-031-98281-1_11","title":"Language Models for Educational Question Generation: Practical Challenges, Personalization Opportunities, and Parameter Optimization","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Computer science; Personalization; Data science; Artificial intelligence; World Wide Web","score_opus":0.09678711961009394,"score_gpt":0.3131462608422016,"score_spread":0.21635914123210764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412684160","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008754825,0.0008930292,0.98227376,0.0016070617,0.000099611825,0.00011508522,0.0004951871,0.0034963973,0.0022650356],"genre_scores_gemma":[0.26430365,0.0012499398,0.7214191,0.00081123196,0.00033485328,0.0005433871,0.0022887553,0.0013972787,0.0076518566],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99738425,0.0016262837,0.00014583385,0.0004390072,0.00027405174,0.00013050501],"domain_scores_gemma":[0.9870569,0.009734035,0.00022361457,0.001690695,0.000988788,0.00030607657],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005601637,0.0008496313,0.0013379883,0.00083477603,0.0005713489,0.0036607457,0.0025077583,0.0017795052,0.009852931],"category_scores_gemma":[0.022383556,0.0008226152,0.0014773457,0.0011206297,0.0006141877,0.0063416255,0.002093154,0.00389888,0.005653994],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000502314,0.0004430661,0.0019392386,0.0004472153,0.00018451962,0.00011187679,0.0007919449,0.15438037,0.008177112,0.051632877,0.029110534,0.7522789],"study_design_scores_gemma":[0.000037203397,0.00004509905,0.00022333537,0.000043843924,0.000043878143,0.00006045732,0.00013922334,0.911026,0.0025133367,0.07692734,0.0089148795,0.000025449828],"about_ca_topic_score_codex":0.0041658804,"about_ca_topic_score_gemma":0.007026898,"teacher_disagreement_score":0.009852931,"about_ca_system_score_codex":0.001425782,"about_ca_system_score_gemma":0.0014747594,"threshold_uncertainty_score":0.03296131},"labels":[],"label_agreement":null},{"id":"W4412699824","doi":"10.3390/robotics14080102","title":"MLLM-Search: A Zero-Shot Approach to Finding People Using Multimodal Large Language Models","year":2025,"lang":"en","type":"article","venue":"Robotics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Health Sciences North; Baycrest Hospital; Toronto Rehabilitation Institute; University of Toronto","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; AGE-WELL","keywords":"Zero (linguistics); Shot (pellet); Computer science; Artificial intelligence; Linguistics; Natural language processing; Philosophy","score_opus":0.0891568815851013,"score_gpt":0.32376948483380474,"score_spread":0.23461260324870345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412699824","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00757626,0.0003576211,0.98538,0.0001890929,0.00004379683,0.00009887384,0.00026014683,0.004657622,0.0014364544],"genre_scores_gemma":[0.27958822,0.0005088594,0.7088145,0.00083706883,0.000073365045,0.0004940365,0.0019040665,0.0010862912,0.0066935504],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99905723,0.00030294122,0.000045483783,0.00030247617,0.00020012826,0.00009181057],"domain_scores_gemma":[0.99894375,0.0006442509,0.000066601635,0.00013111094,0.00013736445,0.000076970515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010813494,0.0016489923,0.001482275,0.0010234563,0.0006474424,0.0014606849,0.0036160494,0.0021898742,0.0053508338],"category_scores_gemma":[0.0040595843,0.0008660817,0.0018717856,0.00082363613,0.00089289626,0.0028834243,0.003257784,0.0018460662,0.0018474725],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073834124,0.00037850146,0.0017369878,0.0008613097,0.00033195427,0.00079427974,0.0018105612,0.39024875,0.02703063,0.024433477,0.016558027,0.5350772],"study_design_scores_gemma":[0.00003128023,0.00011227763,0.0001606998,0.000025228846,0.000030166362,0.000116451796,0.00019796731,0.97778386,0.0036007566,0.0146180345,0.003292609,0.00003073647],"about_ca_topic_score_codex":0.012213746,"about_ca_topic_score_gemma":0.018825443,"teacher_disagreement_score":0.012213746,"about_ca_system_score_codex":0.0010125908,"about_ca_system_score_gemma":0.0015155055,"threshold_uncertainty_score":0.024285316},"labels":[],"label_agreement":null},{"id":"W4412827352","doi":"10.1145/3744340","title":"AcTracer: Active Testing of Large Language Model via Multi-Stage Sampling","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Sampling (signal processing)","score_opus":0.15140591138260573,"score_gpt":0.37414789205765875,"score_spread":0.22274198067505302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412827352","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064208195,0.0013472639,0.9052391,0.0005872447,0.00026543194,0.0006600775,0.0006934571,0.02486438,0.0021348344],"genre_scores_gemma":[0.57811195,0.00038033532,0.4078541,0.0014972511,0.00021184691,0.0016934989,0.003951614,0.002447169,0.0038522603],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99015146,0.005610242,0.000496982,0.0017231375,0.0015311949,0.0004870531],"domain_scores_gemma":[0.96753657,0.02500265,0.00077810895,0.0034258778,0.0024920995,0.0007646354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011749129,0.0035591677,0.0024953613,0.0013515328,0.00099207,0.0022773047,0.006859713,0.0025052489,0.0037476954],"category_scores_gemma":[0.034601208,0.0011206681,0.0021141658,0.00075698306,0.0016149215,0.0060859127,0.004505344,0.0044452786,0.002830525],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025666906,0.0014441628,0.013630121,0.00090796745,0.0008449425,0.0008154339,0.0009951546,0.23938479,0.034418933,0.007808046,0.019908428,0.6772753],"study_design_scores_gemma":[0.0001527707,0.00034321164,0.0004546003,0.000021271033,0.00005514472,0.000110347464,0.00006766679,0.9849727,0.008145188,0.0042804787,0.0013565146,0.00003996337],"about_ca_topic_score_codex":0.005552365,"about_ca_topic_score_gemma":0.007547141,"teacher_disagreement_score":0.011749129,"about_ca_system_score_codex":0.0009395629,"about_ca_system_score_gemma":0.0027433205,"threshold_uncertainty_score":0.062136114},"labels":[],"label_agreement":null},{"id":"W4412841121","doi":"10.1609/aaaiss.v6i1.36064","title":"Creative Thought Embeddings: A Framework for Instilling Creativity in Large Language Models","year":2025,"lang":"en","type":"article","venue":"Proceedings of the AAAI Symposium Series","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Creativity; Cognitive science; Epistemology; Computer science; Psychology; Philosophy; Social psychology","score_opus":0.012895686172021036,"score_gpt":0.2732768155577451,"score_spread":0.26038112938572405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412841121","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046431427,0.00010064672,0.99226123,0.00021693569,0.000025156809,0.0000846607,0.00018037228,0.0015226325,0.0009652047],"genre_scores_gemma":[0.1789367,0.00021755336,0.81746703,0.0001340925,0.000048293183,0.0004795586,0.0007211441,0.00040590245,0.0015896868],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979097,0.0011985374,0.00016166219,0.0003710825,0.000274972,0.000084054045],"domain_scores_gemma":[0.99063665,0.0063093444,0.0006286971,0.0016326893,0.00046020243,0.00033240695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044707693,0.0014698193,0.0006458699,0.0018776801,0.00083313906,0.0033349206,0.0023720078,0.0013998495,0.0057865796],"category_scores_gemma":[0.019671649,0.0009317917,0.0023332094,0.0013488272,0.0024196645,0.006896519,0.004552339,0.0028026775,0.0012075516],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036163745,0.00032759327,0.0043060896,0.0006669529,0.00022819794,0.00035231034,0.0031335652,0.2750166,0.0071983775,0.4191977,0.0050742975,0.2841367],"study_design_scores_gemma":[0.00003575802,0.00007366797,0.00017667547,0.00006266954,0.000030988336,0.000060721653,0.00016175803,0.74176985,0.0017413761,0.24904168,0.006814808,0.000030104411],"about_ca_topic_score_codex":0.0023179296,"about_ca_topic_score_gemma":0.0050055394,"teacher_disagreement_score":0.0057865796,"about_ca_system_score_codex":0.0013376221,"about_ca_system_score_gemma":0.0014596866,"threshold_uncertainty_score":0.02364397},"labels":[],"label_agreement":null},{"id":"W4412876920","doi":"10.1145/3711896.3737233","title":"Hierarchical Lexical Graph for Enhanced Multi-Hop Retrieval","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Graph; Hop (telecommunications); Information retrieval; Artificial intelligence; Theoretical computer science; Computer network","score_opus":0.051062268313731306,"score_gpt":0.32389812752218006,"score_spread":0.27283585920844877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412876920","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022886971,0.0034017162,0.92848754,0.0010282865,0.0001837739,0.0005030636,0.0073374454,0.03144278,0.0047284327],"genre_scores_gemma":[0.27873766,0.0010799449,0.6838491,0.000933011,0.00020861835,0.00056542666,0.025963482,0.001743365,0.0069193505],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985386,0.0005412832,0.00011679598,0.00036231204,0.00033870473,0.00010242566],"domain_scores_gemma":[0.99688643,0.001764647,0.0001604932,0.0007056565,0.0003960738,0.00008659463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016558109,0.0011691427,0.0010804819,0.004641921,0.0007825833,0.001675334,0.0021682335,0.0016473816,0.0084334565],"category_scores_gemma":[0.009576448,0.00046516783,0.0012782776,0.0032831002,0.0007722132,0.004267349,0.0028182743,0.0012298896,0.004632644],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006174087,0.0003156648,0.0029224895,0.0015823798,0.00026449314,0.0006270881,0.0010173195,0.07842958,0.027743043,0.04164293,0.08250236,0.76233524],"study_design_scores_gemma":[0.00020417711,0.0002481703,0.0011789505,0.00012481121,0.00020890722,0.00052749366,0.00047507242,0.7865539,0.015711593,0.14999232,0.04467739,0.00009711735],"about_ca_topic_score_codex":0.0069286735,"about_ca_topic_score_gemma":0.014266315,"teacher_disagreement_score":0.0084334565,"about_ca_system_score_codex":0.0010828112,"about_ca_system_score_gemma":0.0018378227,"threshold_uncertainty_score":0.028212726},"labels":[],"label_agreement":null},{"id":"W4412877066","doi":"10.1145/3711896.3737435","title":"CURE: A dataset for Clinical Understanding &amp; Retrieval Evaluation","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Montreal Clinical Research Institute","funders":"","keywords":"Computer science; Information retrieval; Artificial intelligence","score_opus":0.4613842081861084,"score_gpt":0.4942159585426222,"score_spread":0.032831750356513756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412877066","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.071945384,0.007199789,0.018370671,0.0032650903,0.0007860741,0.0039768713,0.86788756,0.009441182,0.017127302],"genre_scores_gemma":[0.039209347,0.0008007589,0.020302422,0.00058700703,0.0001865235,0.001923299,0.93310297,0.00031991475,0.0035676584],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953507,0.0015002143,0.000896412,0.0007769661,0.0012132317,0.0002623913],"domain_scores_gemma":[0.9900013,0.0040544313,0.0009257989,0.001979047,0.0021955122,0.0008439083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004433578,0.0018656718,0.0014715903,0.0075775925,0.0010480498,0.0019476227,0.0025066207,0.002759043,0.009752259],"category_scores_gemma":[0.021277757,0.0003608287,0.0016553428,0.00457292,0.00078120077,0.0019476403,0.002307844,0.0018071927,0.012311399],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012123253,0.0010205357,0.012096189,0.005186322,0.00046453104,0.0006761688,0.000529825,0.004040884,0.0061065154,0.0019831252,0.83182496,0.13485864],"study_design_scores_gemma":[0.002650017,0.0024075836,0.09097033,0.0015256231,0.0007764244,0.005815408,0.0018751485,0.043290354,0.01780334,0.0063467715,0.8259997,0.00053937314],"about_ca_topic_score_codex":0.010079062,"about_ca_topic_score_gemma":0.017738277,"teacher_disagreement_score":0.010079062,"about_ca_system_score_codex":0.0018354004,"about_ca_system_score_gemma":0.002663004,"threshold_uncertainty_score":0.032624483},"labels":[],"label_agreement":null},{"id":"W4412886529","doi":"10.21203/rs.3.rs-7226688/v1","title":"Large Language Models as Mediators: Addressing Rater Disagreement in Turkish Essay Scoring","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Rubric; Turkish; Writing assessment; Task (project management); Test (biology); Active listening; Psychology; Consistency (knowledge bases); Benchmark (surveying); Applied psychology; Computer science; Cognitive psychology; Mathematics education; Linguistics; Artificial intelligence","score_opus":0.09587885843140762,"score_gpt":0.42052009158985515,"score_spread":0.3246412331584475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412886529","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60476446,0.0029743637,0.37494224,0.003915283,0.00045497302,0.00055882987,0.00072851207,0.001345337,0.010316017],"genre_scores_gemma":[0.9750676,0.0001158205,0.02258132,0.0001307849,0.0001422325,0.00029666,0.0004597208,0.00024841394,0.0009574211],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.83612156,0.14691986,0.003304284,0.007274671,0.004659792,0.0017197746],"domain_scores_gemma":[0.588212,0.3664362,0.01128105,0.017561823,0.014258907,0.002250006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10755869,0.0018595116,0.0020249612,0.0026361998,0.0030917665,0.007873283,0.003719653,0.0034538405,0.0036330684],"category_scores_gemma":[0.3484689,0.0010870296,0.0014385032,0.0030821506,0.0020876036,0.008003675,0.0070511145,0.0047362945,0.0018123724],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0107624745,0.0016262712,0.27240148,0.0017934319,0.0043089124,0.0013897142,0.036986224,0.06064055,0.015731396,0.06103264,0.017549125,0.51577777],"study_design_scores_gemma":[0.00060138415,0.0008702424,0.07047709,0.00043476265,0.002761425,0.00072273344,0.0069078067,0.8139919,0.019097453,0.0734902,0.010313517,0.0003315243],"about_ca_topic_score_codex":0.0026221268,"about_ca_topic_score_gemma":0.003219574,"teacher_disagreement_score":0.10755869,"about_ca_system_score_codex":0.002134817,"about_ca_system_score_gemma":0.0029620754,"threshold_uncertainty_score":0.56883156},"labels":[],"label_agreement":null},{"id":"W4412887682","doi":"10.18653/v1/2025.findings-acl.1134","title":"Small Encoders Can Rival Large Decoders in Detecting Groundedness","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Encoder; Computer science; Decoding methods; Algorithm","score_opus":0.028727232356802322,"score_gpt":0.25718900277262124,"score_spread":0.22846177041581892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412887682","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09585993,0.004063456,0.84938866,0.0018299759,0.00027474665,0.00038408986,0.0019575977,0.041508134,0.0047334847],"genre_scores_gemma":[0.58700496,0.0013374089,0.39948747,0.00093133724,0.0001939633,0.0004004929,0.004531214,0.0016880864,0.004425092],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99789065,0.0008497755,0.00014391675,0.0005588284,0.0004021556,0.00015465154],"domain_scores_gemma":[0.9841144,0.012003415,0.0003833188,0.0022855944,0.0008858016,0.00032744202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005027584,0.0017693983,0.0010937576,0.0013569716,0.00061374676,0.0023235306,0.002202496,0.0015744211,0.0053238887],"category_scores_gemma":[0.024307854,0.0011863128,0.0010098668,0.0010633573,0.0011562459,0.007238635,0.003222714,0.003420665,0.0037932033],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015131843,0.000469529,0.0110766655,0.0012417851,0.0004536865,0.0003890621,0.0009272963,0.064401284,0.053725995,0.017594423,0.023154076,0.82505304],"study_design_scores_gemma":[0.00012229085,0.00022384814,0.0016908043,0.000098128985,0.00022148808,0.00021300076,0.00023206242,0.92584896,0.02586378,0.036493145,0.008933618,0.000058900725],"about_ca_topic_score_codex":0.007939496,"about_ca_topic_score_gemma":0.019271115,"teacher_disagreement_score":0.007939496,"about_ca_system_score_codex":0.0014517433,"about_ca_system_score_gemma":0.002252158,"threshold_uncertainty_score":0.026588678},"labels":[],"label_agreement":null},{"id":"W4412887694","doi":"10.18653/v1/2025.findings-acl.1170","title":"Q-STRUM Debate: Query-Driven Contrastive Summarization for Recommendation Comparison","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea","keywords":"Automatic summarization; Computer science; Information retrieval; Natural language processing","score_opus":0.0249753236133718,"score_gpt":0.30381596742988365,"score_spread":0.27884064381651186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412887694","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021181019,0.0033376764,0.9533262,0.0009293787,0.0003417438,0.00077115605,0.00479218,0.01296631,0.0023543057],"genre_scores_gemma":[0.16344738,0.0006477172,0.81501496,0.00063955784,0.00054043805,0.00095800735,0.014558644,0.0007679328,0.0034254177],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99425054,0.002534177,0.00052793033,0.0014260946,0.001062762,0.00019854582],"domain_scores_gemma":[0.9860524,0.008779435,0.00081493234,0.0018304137,0.0021924386,0.00033035057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006856274,0.0019157823,0.0017865475,0.0052035465,0.0012286309,0.002487651,0.0029288603,0.002327553,0.0068200617],"category_scores_gemma":[0.035425432,0.0005287414,0.0016655434,0.0039461358,0.0007673593,0.0051515913,0.002521628,0.0025149302,0.0031838317],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014302877,0.00039889707,0.0039805323,0.0016465957,0.0004956752,0.00021775988,0.0015504357,0.03573756,0.020702766,0.015723553,0.04961898,0.86849695],"study_design_scores_gemma":[0.0005848781,0.0010032285,0.0031388737,0.0001924346,0.00032627527,0.00028264025,0.00092421216,0.8558476,0.023121975,0.059828363,0.05455356,0.0001961219],"about_ca_topic_score_codex":0.0037608552,"about_ca_topic_score_gemma":0.0077067283,"teacher_disagreement_score":0.006856274,"about_ca_system_score_codex":0.0013992578,"about_ca_system_score_gemma":0.0022214572,"threshold_uncertainty_score":0.03625989},"labels":[],"label_agreement":null},{"id":"W4412888277","doi":"10.18653/v1/2025.findings-acl.695","title":"When Detection Fails: The Power of Fine-Tuned Models to Generate Human-Like Social Media Text","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"Ministère de la Défense Nationale","keywords":"Social media; Computer science; Power (physics); Data science; World Wide Web; Physics","score_opus":0.04444604225395427,"score_gpt":0.2662851425482298,"score_spread":0.22183910029427556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412888277","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44754806,0.0018889843,0.5182476,0.004373233,0.00049991114,0.0006013654,0.0018542041,0.017371183,0.007615388],"genre_scores_gemma":[0.9305428,0.00013429185,0.064831264,0.0009314984,0.0001442603,0.0001449548,0.0015819022,0.00053083815,0.0011582419],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9888831,0.0053438065,0.0004813916,0.003227004,0.0015712769,0.0004934141],"domain_scores_gemma":[0.91600573,0.0634865,0.0034644096,0.014132223,0.0019872945,0.0009239861],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013019296,0.0017404824,0.0012739883,0.0017284538,0.0011993296,0.0031993643,0.0018966754,0.0032722794,0.0013764598],"category_scores_gemma":[0.10084632,0.0006836154,0.0009630917,0.0007749917,0.0021826385,0.0056629903,0.002455344,0.004455468,0.0017070082],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026318438,0.0011296923,0.17217553,0.001633779,0.0008890062,0.0019776458,0.005663967,0.29910496,0.047420572,0.027704263,0.050336704,0.38933212],"study_design_scores_gemma":[0.00009600273,0.00025390895,0.007121521,0.00011038177,0.00005382431,0.00065772067,0.00047708573,0.9470005,0.013028572,0.025057334,0.006075814,0.00006737278],"about_ca_topic_score_codex":0.0036289885,"about_ca_topic_score_gemma":0.0035802869,"teacher_disagreement_score":0.013019296,"about_ca_system_score_codex":0.0010925689,"about_ca_system_score_gemma":0.0011023013,"threshold_uncertainty_score":0.06885344},"labels":[],"label_agreement":null},{"id":"W4412888700","doi":"10.18653/v1/2025.findings-acl.235","title":"R3Mem: Bridging Memory Retention and Retrieval via Reversible Compression","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Bridging (networking); Computer science; Compression (physics); Materials science; Composite material; Computer network","score_opus":0.01827407981595656,"score_gpt":0.24670978942283095,"score_spread":0.2284357096068744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412888700","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08748907,0.0015861007,0.89306796,0.00039722896,0.00014132359,0.0001546182,0.00033690513,0.009416709,0.0074099936],"genre_scores_gemma":[0.82688487,0.0006807806,0.1640226,0.00032328122,0.000057083642,0.00027222853,0.0005334537,0.00045573505,0.0067700674],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99972767,0.00004843688,0.000026281887,0.00006939672,0.00008247937,0.00004563792],"domain_scores_gemma":[0.999297,0.00021873179,0.000059960636,0.00029151715,0.00009895431,0.000033757882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00056839274,0.0006427348,0.0005329194,0.00059752754,0.00041221592,0.0010096023,0.0026463305,0.000591703,0.0047468706],"category_scores_gemma":[0.0025433335,0.0002631071,0.00038301453,0.000605363,0.00072472944,0.0033996177,0.0020160002,0.0007514059,0.0012440496],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009419953,0.00025335743,0.0024180503,0.000548037,0.00009606455,0.0005331842,0.00053115084,0.14525324,0.1039362,0.049765285,0.011039288,0.6846842],"study_design_scores_gemma":[0.000064116604,0.00031064716,0.0003092449,0.000044681492,0.00006752098,0.00030641668,0.00012976292,0.84747505,0.11060285,0.025869114,0.014766593,0.000054030177],"about_ca_topic_score_codex":0.0024040032,"about_ca_topic_score_gemma":0.0032689294,"teacher_disagreement_score":0.0047468706,"about_ca_system_score_codex":0.0006923085,"about_ca_system_score_gemma":0.00090449065,"threshold_uncertainty_score":0.015879929},"labels":[],"label_agreement":null},{"id":"W4412888901","doi":"10.18653/v1/2025.fever-1.19","title":"SANCTUARY: An Efficient Evidence-based Automated Fact Checking System","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science","score_opus":0.043862993180243316,"score_gpt":0.2970264958567863,"score_spread":0.253163502676543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412888901","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039899014,0.003015946,0.63380176,0.0043562544,0.00073921157,0.0010018783,0.02694988,0.27634004,0.013896013],"genre_scores_gemma":[0.22708704,0.0010915224,0.7012688,0.0014084484,0.00043996176,0.00042633855,0.0562828,0.00369849,0.008296642],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99707234,0.0007171566,0.00029514646,0.0006862744,0.0010930428,0.0001360496],"domain_scores_gemma":[0.9846503,0.009361329,0.0009678119,0.0021217584,0.0025306423,0.00036828668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034754302,0.0011659386,0.0010573862,0.0053389217,0.0011680645,0.0034583367,0.0026591732,0.0018609148,0.012505695],"category_scores_gemma":[0.032832135,0.0006955429,0.0014088652,0.0017591253,0.0007210373,0.006754126,0.0043824883,0.0018458904,0.0072768354],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014201993,0.00038705842,0.007932351,0.0016532812,0.00031717183,0.002150062,0.0010960157,0.018630933,0.017531786,0.034393154,0.2687535,0.64573455],"study_design_scores_gemma":[0.00047724767,0.0002864879,0.003363844,0.00047941998,0.00034041193,0.001643236,0.0006397154,0.7030589,0.035544816,0.07011306,0.18381454,0.00023827069],"about_ca_topic_score_codex":0.007381315,"about_ca_topic_score_gemma":0.010757959,"teacher_disagreement_score":0.012505695,"about_ca_system_score_codex":0.0011938802,"about_ca_system_score_gemma":0.004119056,"threshold_uncertainty_score":0.041835666},"labels":[],"label_agreement":null},{"id":"W4412888945","doi":"10.18653/v1/2025.findings-acl.54","title":"Can Large Language Models Address Open-Target Stance Detection?","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Computer science; Artificial intelligence; Natural language processing","score_opus":0.019637186148660408,"score_gpt":0.28532996286241996,"score_spread":0.26569277671375957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412888945","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17132245,0.008959916,0.7534074,0.007325609,0.001270201,0.0005479098,0.007216075,0.036216393,0.013733972],"genre_scores_gemma":[0.7250785,0.0017057327,0.25446245,0.0014295338,0.0005331573,0.00041836782,0.010302916,0.0019309781,0.0041383873],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99721587,0.001629171,0.00017551222,0.0005527226,0.00027332766,0.00015342268],"domain_scores_gemma":[0.9847151,0.011647479,0.00053425325,0.001603019,0.0011069131,0.00039312994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051336093,0.0020925356,0.0011727699,0.0013391944,0.0005950763,0.0028073685,0.0018036197,0.0020501092,0.0048757577],"category_scores_gemma":[0.024774434,0.00071122864,0.0011689266,0.0010514255,0.0006383571,0.0072664777,0.0017019333,0.002679869,0.007637224],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020399746,0.0008008511,0.017394962,0.0023135764,0.00073147396,0.00069520716,0.0014449472,0.1637051,0.03166307,0.013631443,0.068342105,0.6972373],"study_design_scores_gemma":[0.0001599945,0.00025081195,0.0012705756,0.0001281605,0.000118659606,0.00026060102,0.00033113552,0.9510793,0.006990769,0.025213087,0.014145343,0.00005166609],"about_ca_topic_score_codex":0.0035380595,"about_ca_topic_score_gemma":0.0070699044,"teacher_disagreement_score":0.0051336093,"about_ca_system_score_codex":0.0008574341,"about_ca_system_score_gemma":0.0013985537,"threshold_uncertainty_score":0.027149498},"labels":[],"label_agreement":null},{"id":"W4412889420","doi":"10.18653/v1/2025.acl-short.97","title":"Counterfactual-Consistency Prompting for Relative Temporal Understanding in Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; National Research Foundation of Korea; Seoul National University; National Research Foundation","keywords":"Counterfactual thinking; Computer science; Consistency (knowledge bases); Artificial intelligence; Natural language processing; Psychology; Social psychology","score_opus":0.0692207462450151,"score_gpt":0.3069790829925247,"score_spread":0.23775833674750957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412889420","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011108854,0.00026447576,0.98087305,0.00080767466,0.000069866626,0.00014737781,0.0005693363,0.005216441,0.00094282586],"genre_scores_gemma":[0.30092427,0.00018152967,0.69410515,0.0005455541,0.00016322774,0.000303244,0.0020161222,0.00082726934,0.000933593],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9907736,0.0055126552,0.00053648173,0.0018499858,0.0011090281,0.00021823608],"domain_scores_gemma":[0.94357914,0.0435303,0.003131064,0.0072217532,0.0017853727,0.00075243635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013402568,0.0014708472,0.0012485464,0.0017365722,0.0011796672,0.0033240688,0.0031697403,0.002333901,0.006362199],"category_scores_gemma":[0.07180468,0.00091838255,0.0018979757,0.0013305455,0.0016957717,0.009394937,0.0057285,0.0054006567,0.0013687955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015665009,0.00051712996,0.0097428635,0.0015457919,0.00038827103,0.0010776434,0.007563139,0.17107427,0.017897349,0.24402168,0.021919351,0.52268595],"study_design_scores_gemma":[0.00010838615,0.000062164465,0.00040029114,0.00009395716,0.00007708896,0.0001943519,0.00035818224,0.7707428,0.0060543697,0.21083984,0.011010812,0.000057698326],"about_ca_topic_score_codex":0.0027389263,"about_ca_topic_score_gemma":0.005766189,"teacher_disagreement_score":0.013402568,"about_ca_system_score_codex":0.0017687692,"about_ca_system_score_gemma":0.0035140188,"threshold_uncertainty_score":0.07088041},"labels":[],"label_agreement":null},{"id":"W4412906974","doi":"10.3390/fire8080306","title":"Entity Recognition Method for Fire Safety Standards Based on FT-FLAT","year":2025,"lang":"en","type":"article","venue":"Fire","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"123 Certification (Canada)","funders":"","keywords":"Fire safety; Computer science; Forensic engineering; Environmental science; Business; Engineering; Risk analysis (engineering)","score_opus":0.028265250961732998,"score_gpt":0.32030161992409734,"score_spread":0.29203636896236435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412906974","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058258668,0.0022963241,0.8958418,0.0007615198,0.00037540428,0.00048256313,0.009573836,0.01818349,0.014226452],"genre_scores_gemma":[0.49106154,0.0017329192,0.4317115,0.00042903746,0.00024121035,0.00037607277,0.055906497,0.00059087586,0.017950257],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990816,0.00007412678,0.000114554416,0.00040046903,0.00023161958,0.000097519434],"domain_scores_gemma":[0.99927944,0.00015110655,0.00009003936,0.00017401215,0.00026779514,0.000037615904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00064905424,0.00097500713,0.00060799764,0.004429095,0.00064921635,0.001127827,0.0015297468,0.001042164,0.0047353082],"category_scores_gemma":[0.0023589064,0.0002618733,0.0014354918,0.0033809687,0.00031260357,0.0037861627,0.0011814743,0.0011482167,0.0036242038],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025054757,0.00018525842,0.006897153,0.00041533235,0.00017316452,0.0006712781,0.00029105524,0.036808833,0.023025392,0.018433489,0.05047068,0.86237776],"study_design_scores_gemma":[0.000033154218,0.0000911452,0.0056779045,0.000068839116,0.00016583277,0.00062715,0.00025189866,0.8943688,0.03094867,0.020492451,0.047208484,0.00006573337],"about_ca_topic_score_codex":0.014381411,"about_ca_topic_score_gemma":0.017058216,"teacher_disagreement_score":0.014381411,"about_ca_system_score_codex":0.0011176548,"about_ca_system_score_gemma":0.0014929965,"threshold_uncertainty_score":0.028595388},"labels":[],"label_agreement":null},{"id":"W4412944217","doi":"10.18653/v1/2025.trl-1.14","title":"OrQA – Open Data Retrieval for Question Answering dataset generation","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Question answering; Information retrieval","score_opus":0.16214780063289805,"score_gpt":0.39911111832994134,"score_spread":0.2369633176970433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412944217","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022418978,0.0019229084,0.38602084,0.0022214714,0.00076040346,0.00301484,0.357056,0.21118361,0.015400977],"genre_scores_gemma":[0.04757511,0.00035086108,0.4406008,0.0007979874,0.00009882829,0.0023085962,0.5005195,0.0041255364,0.0036228178],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961116,0.0011304172,0.0004208353,0.0011594839,0.0009417947,0.00023592856],"domain_scores_gemma":[0.9940943,0.0023113657,0.00026031816,0.0020000895,0.0010483622,0.00028546006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004352254,0.001973546,0.0008710693,0.0038932872,0.0010459044,0.0026275858,0.0036616516,0.0016798194,0.016730884],"category_scores_gemma":[0.018430764,0.000599601,0.002229046,0.003028749,0.0008285468,0.0034165662,0.0040618307,0.0026304754,0.01098701],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007760475,0.00069108093,0.006713685,0.0031983214,0.00029525626,0.00048032668,0.0016091699,0.02295367,0.016414838,0.028764103,0.7061117,0.21199177],"study_design_scores_gemma":[0.00060124253,0.0003443366,0.004475196,0.00032336355,0.00010259801,0.00044778787,0.0011849159,0.3316324,0.031436473,0.06915513,0.5601192,0.00017737069],"about_ca_topic_score_codex":0.011369805,"about_ca_topic_score_gemma":0.016495012,"teacher_disagreement_score":0.016730884,"about_ca_system_score_codex":0.0015235426,"about_ca_system_score_gemma":0.0025425642,"threshold_uncertainty_score":0.05597037},"labels":[],"label_agreement":null},{"id":"W4412980931","doi":"10.1016/j.datak.2025.102499","title":"Corrigendum to “Large Language Models for Conceptual Modeling: Assessment and Application Potential” [Knowledge and Data Engineering Volume 160, November 2025, 102480]","year":2025,"lang":"en","type":"erratum","venue":"Data & Knowledge Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Volume (thermodynamics); Computer science; Conceptual model; Data science; Database; Physics","score_opus":0.04627239330600777,"score_gpt":0.3163658736183723,"score_spread":0.2700934803123646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412980931","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00022045914,0.0017668651,0.004826576,0.10807558,0.83411396,0.00013338917,0.004434102,0.0020546748,0.0443744],"genre_scores_gemma":[0.0040309755,0.0035380505,0.0068503157,0.063892834,0.09925779,0.00030240865,0.008664257,0.0022690278,0.8111943],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9929202,0.0011700835,0.00068645034,0.0008516067,0.00386456,0.00050719414],"domain_scores_gemma":[0.9663382,0.0053836,0.000722807,0.001839002,0.02476387,0.0009524017],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049090427,0.0024690328,0.0023804677,0.006057404,0.0059609218,0.0088021895,0.0037076895,0.0079813665,0.21374533],"category_scores_gemma":[0.041025594,0.0013573879,0.002527323,0.004284673,0.0017751093,0.0045255115,0.0029044582,0.007878931,0.1587043],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007958095,0.000003677651,0.000011350639,0.000018875699,0.0000022005268,0.000014127528,0.0000038442413,0.00002608441,0.000017864711,0.0005266033,0.9973673,0.002000152],"study_design_scores_gemma":[0.000024976105,0.00001932901,0.00046932453,0.00015210902,0.000026757885,0.00005684229,0.000050604478,0.00080986705,0.00028241175,0.0036418273,0.9944285,0.000037373837],"about_ca_topic_score_codex":0.08813502,"about_ca_topic_score_gemma":0.13421507,"teacher_disagreement_score":0.21374533,"about_ca_system_score_codex":0.008395775,"about_ca_system_score_gemma":0.006837992,"threshold_uncertainty_score":0.7150494},"labels":[],"label_agreement":null},{"id":"W4413002442","doi":"10.3389/frai.2025.1592013","title":"Large language models for closed-library multi-document query, test generation, and evaluation","year":2025,"lang":"en","type":"article","venue":"Frontiers in Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Massachusetts Institute of Technology","keywords":"Computer science; Knowledge base; Leverage (statistics); Set (abstract data type); Information retrieval; World Wide Web; Language model; Data science; Artificial intelligence","score_opus":0.05380820147992983,"score_gpt":0.32675085891863037,"score_spread":0.27294265743870055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413002442","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021052776,0.0024660141,0.9462829,0.0012949394,0.0002237139,0.0014113192,0.004442095,0.017722586,0.0051036663],"genre_scores_gemma":[0.3579359,0.0007049435,0.6134058,0.0009590129,0.00028966408,0.003387264,0.015857896,0.0014731688,0.0059863506],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98699385,0.007724641,0.0009390887,0.0016296916,0.0022548202,0.00045795078],"domain_scores_gemma":[0.9630075,0.027635114,0.0013463072,0.0033725877,0.003958421,0.00068016833],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015154137,0.0019536624,0.0017019372,0.0023564803,0.0007875046,0.0033508807,0.003948725,0.003371797,0.01296897],"category_scores_gemma":[0.049891915,0.0007989362,0.0021538844,0.0021318875,0.0012627238,0.0036291333,0.0028957166,0.0030455682,0.006455894],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014918479,0.0009247357,0.0036454422,0.0011583087,0.00033699512,0.00033554272,0.00036248413,0.4165642,0.0032946956,0.025922634,0.048920803,0.49704233],"study_design_scores_gemma":[0.000089669105,0.000116360825,0.00039604903,0.00005252792,0.000029603621,0.00007919363,0.00004292588,0.97593224,0.0013623446,0.01822986,0.0036438645,0.000025430656],"about_ca_topic_score_codex":0.011863116,"about_ca_topic_score_gemma":0.010965212,"teacher_disagreement_score":0.015154137,"about_ca_system_score_codex":0.0040592393,"about_ca_system_score_gemma":0.0029292488,"threshold_uncertainty_score":0.08014369},"labels":[],"label_agreement":null},{"id":"W4413025253","doi":"10.1007/978-3-032-00891-6_28","title":"Enhancing Ultra-Low-Bit Quantization of Large Language Models Through Saliency-Aware Partial Retraining","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University of New Brunswick","funders":"","keywords":"Computer science; Retraining; Quantization (signal processing); Bit (key); Artificial intelligence; Speech recognition; Computer vision; Computer network","score_opus":0.02045388212035081,"score_gpt":0.2747283330984558,"score_spread":0.25427445097810497,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413025253","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040352635,0.0022657162,0.9460774,0.00075821526,0.00038216356,0.00007545922,0.0005043741,0.006314085,0.0032699015],"genre_scores_gemma":[0.5857134,0.0011297582,0.39934537,0.0010517242,0.0002804849,0.00013198976,0.0014829319,0.00068150193,0.0101828175],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9996043,0.000085941545,0.0000252306,0.00009064831,0.00014402371,0.000049781003],"domain_scores_gemma":[0.99896145,0.00055221544,0.000055563407,0.00016437213,0.00020958275,0.000056781064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00062067545,0.0010579742,0.0008362975,0.00045115987,0.00036355958,0.00081889366,0.0011353915,0.00082155335,0.0055093006],"category_scores_gemma":[0.0034905758,0.0003350181,0.0004119924,0.0006101517,0.0003931438,0.0018981753,0.001258148,0.0015181133,0.0022987658],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054365495,0.00025107694,0.0006342722,0.000254359,0.00007527662,0.00022458835,0.00017394345,0.067050144,0.09183444,0.00594007,0.011309206,0.82170886],"study_design_scores_gemma":[0.000023386614,0.000115176954,0.00035165536,0.00002014501,0.000026801226,0.000100698926,0.000031882653,0.97298265,0.018088335,0.0058856127,0.0023545583,0.000019023573],"about_ca_topic_score_codex":0.0058100577,"about_ca_topic_score_gemma":0.014151216,"teacher_disagreement_score":0.0058100577,"about_ca_system_score_codex":0.0004497984,"about_ca_system_score_gemma":0.0009116894,"threshold_uncertainty_score":0.018430412},"labels":[],"label_agreement":null},{"id":"W4413029022","doi":"10.1007/s43681-025-00806-5","title":"The deep illusion: a critical analysis of DeepSeek and the limits of large language models","year":2025,"lang":"en","type":"article","venue":"AI and Ethics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cegep de Sainte Foy","funders":"","keywords":"Illusion; Computer science; Linguistics; Psychology; Cognitive psychology; Philosophy","score_opus":0.03481016132546281,"score_gpt":0.3463852527258339,"score_spread":0.3115750914003711,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413029022","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05194015,0.01861645,0.46747276,0.31767738,0.0012545565,0.00012567543,0.00047525906,0.00045015977,0.1419877],"genre_scores_gemma":[0.9387721,0.0034819883,0.035504393,0.012771767,0.0025714696,0.00024990836,0.000091774615,0.00045183412,0.006104766],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.987295,0.0088844,0.00036511556,0.0011441772,0.0018759227,0.00043540596],"domain_scores_gemma":[0.819337,0.16537197,0.0027392888,0.0076784743,0.003667117,0.0012061753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024457788,0.0007534674,0.0014787504,0.0035833027,0.0038545763,0.009105838,0.002988943,0.0055176457,0.006973223],"category_scores_gemma":[0.10903154,0.0009456682,0.000941129,0.0019721547,0.044810764,0.03720691,0.006915932,0.0133159235,0.0005626016],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015557647,0.0000041811627,0.000095023715,0.0000260189,0.0000083214545,0.00001992381,0.0010504405,0.0005459879,0.000033147673,0.9943598,0.0016073852,0.0022341812],"study_design_scores_gemma":[0.0000042202914,0.0000021012627,0.000035585508,0.000020795382,0.0000024218923,0.000014704692,0.00017000842,0.0021358312,0.000040392406,0.9947792,0.0027890503,0.000005658345],"about_ca_topic_score_codex":0.00335317,"about_ca_topic_score_gemma":0.0021343252,"teacher_disagreement_score":0.024457788,"about_ca_system_score_codex":0.0052247383,"about_ca_system_score_gemma":0.0024518627,"threshold_uncertainty_score":0.12934667},"labels":[],"label_agreement":null},{"id":"W4413040066","doi":"10.3233/shti250922","title":"Human in the Loop: Embedding Medical Expert Input in Large Language Models for Clinical Applications","year":2025,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Automatic summarization; Ontology; Dravet syndrome; Epilepsy; Unified Medical Language System; Artificial intelligence; Natural language processing; Data science; Medicine; Psychiatry","score_opus":0.08934316528721445,"score_gpt":0.48167036378910366,"score_spread":0.39232719850188924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413040066","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014153626,0.00036589062,0.9751917,0.0009815253,0.000053676136,0.00018423895,0.0004900182,0.007246773,0.001332613],"genre_scores_gemma":[0.43048492,0.0004897856,0.56298697,0.00087777025,0.00011936358,0.00050899544,0.0018174106,0.0006382876,0.0020765043],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971347,0.001889558,0.00011879156,0.00039365533,0.00036802999,0.00009534162],"domain_scores_gemma":[0.98648804,0.011465755,0.0003826353,0.0008845194,0.0005759351,0.00020317307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005355183,0.0010409495,0.00067346287,0.0012331852,0.0004369646,0.0017131218,0.0014222872,0.001518631,0.0039242697],"category_scores_gemma":[0.024432236,0.0005284834,0.0010126502,0.0009446771,0.0007861348,0.0030232568,0.002311238,0.0014474405,0.0014937943],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012512763,0.00053300057,0.005140395,0.00049267005,0.0002900755,0.0006219558,0.001621535,0.3758378,0.008179623,0.016706044,0.012162659,0.57716286],"study_design_scores_gemma":[0.000038586517,0.00005604804,0.0001623359,0.000023481214,0.000030268078,0.00004680114,0.000074484145,0.97799987,0.0022941716,0.016448053,0.0028101099,0.000015787831],"about_ca_topic_score_codex":0.007836538,"about_ca_topic_score_gemma":0.012487311,"teacher_disagreement_score":0.007836538,"about_ca_system_score_codex":0.001121568,"about_ca_system_score_gemma":0.0020819115,"threshold_uncertainty_score":0.028321207},"labels":[],"label_agreement":null},{"id":"W4413060628","doi":"10.1145/3748316","title":"Inner-character and Inner-word Features Based Representation Learning for Chinese Word Embedding","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Asian and Low-Resource Language Information Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Fundamental Research Funds for the Central Universities; National Key Research and Development Program of China; Natural Science Foundation of Sichuan Province; China Postdoctoral Science Foundation; National Natural Science Foundation of China","keywords":"Pinyin; Computer science; Artificial intelligence; Natural language processing; Word (group theory); Character (mathematics); Word embedding; Feature (linguistics); Similarity (geometry); Chinese characters; Speech recognition; Embedding; Linguistics; Mathematics","score_opus":0.006695338832000711,"score_gpt":0.27081476975019714,"score_spread":0.26411943091819645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413060628","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09279989,0.0012683265,0.90074706,0.00031222307,0.00015617095,0.000089067915,0.0006776812,0.0019658655,0.0019837713],"genre_scores_gemma":[0.8096068,0.001291546,0.17573471,0.00022737183,0.00018881902,0.00025797426,0.00477657,0.00019849358,0.0077176844],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99955684,0.00009982144,0.000039052855,0.0001703793,0.0000793573,0.000054442473],"domain_scores_gemma":[0.99952054,0.00014328968,0.000047187077,0.00010358348,0.0001532299,0.0000321297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004026518,0.0011884577,0.0007506169,0.00091390294,0.0003005513,0.0005149595,0.00083741394,0.00048326945,0.0018379075],"category_scores_gemma":[0.0015443653,0.00026396816,0.00077052996,0.001630798,0.00040574413,0.0023822722,0.0009714679,0.0011489562,0.00090944796],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001711837,0.0002080612,0.0034349652,0.00021261047,0.00012064775,0.000108234104,0.00020066876,0.074397966,0.014162626,0.010067125,0.010495007,0.8864209],"study_design_scores_gemma":[0.000014343927,0.00009720097,0.0008987217,0.000011113808,0.000034263314,0.000057902696,0.000050440958,0.98532087,0.0038984297,0.007786679,0.0018100444,0.000019997644],"about_ca_topic_score_codex":0.0037890018,"about_ca_topic_score_gemma":0.006093041,"teacher_disagreement_score":0.0037890018,"about_ca_system_score_codex":0.00041508404,"about_ca_system_score_gemma":0.0008894226,"threshold_uncertainty_score":0.007533908},"labels":[],"label_agreement":null},{"id":"W4413062930","doi":"10.1145/3715335.3737684","title":"Community-Driven Data Practices for Advancing Ethical and Equitable AI in Low-Resource Language Contexts","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Resource (disambiguation); Knowledge management; Data science","score_opus":0.05674212384201326,"score_gpt":0.3930905079977772,"score_spread":0.33634838415576396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413062930","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066131875,0.00038295897,0.9510857,0.01963408,0.00025807903,0.0006726115,0.00042503484,0.0015967967,0.019331636],"genre_scores_gemma":[0.18072979,0.00046473692,0.8049099,0.0028293617,0.0003364492,0.0021306966,0.0015977487,0.0012518935,0.0057495157],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.8663457,0.106550336,0.0038677368,0.007857916,0.013585477,0.0017928145],"domain_scores_gemma":[0.5992862,0.24339478,0.008322192,0.092650436,0.044759694,0.011586779],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.115542404,0.0010541541,0.0012477095,0.008233692,0.007588975,0.017265214,0.008215643,0.0053722397,0.012260086],"category_scores_gemma":[0.28128672,0.0013309835,0.0014761082,0.007065956,0.012022174,0.030185021,0.024027137,0.0111598205,0.005333862],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017652108,0.0006276876,0.006993663,0.0006969779,0.00017053785,0.00025818587,0.034046162,0.00880065,0.0032011624,0.73080295,0.033219136,0.1810063],"study_design_scores_gemma":[0.000054554384,0.0000497367,0.0006577569,0.00037968738,0.000032216078,0.00015374385,0.0083567165,0.031600043,0.0024339282,0.8363926,0.119816005,0.00007300837],"about_ca_topic_score_codex":0.0072079413,"about_ca_topic_score_gemma":0.011997726,"teacher_disagreement_score":0.115542404,"about_ca_system_score_codex":0.0056477725,"about_ca_system_score_gemma":0.021628063,"threshold_uncertainty_score":0.61105394},"labels":[],"label_agreement":null},{"id":"W4413065313","doi":"10.1007/978-981-96-6294-4_20","title":"Word Embedding Bias in Large Language Models","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Word (group theory); Word embedding; Computer science; Natural language processing; Embedding; Linguistics; Artificial intelligence; Philosophy","score_opus":0.0730375053919994,"score_gpt":0.33158081651414867,"score_spread":0.2585433111221493,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413065313","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02413854,0.008436072,0.9515447,0.003057497,0.0006693206,0.000032090702,0.00058106903,0.0015236252,0.010016971],"genre_scores_gemma":[0.6497983,0.013129883,0.2804224,0.0017890508,0.0035403788,0.00032154538,0.0033165123,0.0025290707,0.045152813],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837387,0.0009914857,0.000072522416,0.00024704874,0.00025099996,0.00006408778],"domain_scores_gemma":[0.97733194,0.020463727,0.00036536678,0.0010248714,0.00067005795,0.00014405028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004371388,0.0007855099,0.0012791996,0.0009869726,0.00048420386,0.0022438583,0.0011198347,0.001134525,0.005111641],"category_scores_gemma":[0.026793983,0.00083250855,0.0006764934,0.0019843583,0.000976442,0.005551112,0.0015373125,0.0030823878,0.0019714546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027258706,0.00010824223,0.003688968,0.0007704298,0.0002982696,0.00029192716,0.0005458983,0.08802928,0.0048146583,0.46825674,0.03848936,0.39443356],"study_design_scores_gemma":[0.000018398516,0.000032029093,0.00070883613,0.00006816621,0.00006253017,0.00013667735,0.000044570064,0.4200761,0.0016022698,0.5698555,0.0073695662,0.000025346284],"about_ca_topic_score_codex":0.0013636813,"about_ca_topic_score_gemma":0.0023726602,"teacher_disagreement_score":0.005111641,"about_ca_system_score_codex":0.00085534016,"about_ca_system_score_gemma":0.0006677528,"threshold_uncertainty_score":0.023118377},"labels":[],"label_agreement":null},{"id":"W4413066139","doi":"10.1145/3733155.3734900","title":"Generalized Follow-Up WHODAS 2.0 Assessment Through Language Models and Adaptive Clustering Ensemble in Higher Education","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; Carleton University","funders":"","keywords":"Cluster analysis; Computer science; Natural language processing; Artificial intelligence; Data science","score_opus":0.06027243842198861,"score_gpt":0.3330583787772414,"score_spread":0.27278594035525283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413066139","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7329907,0.00012405608,0.2607901,0.0003761861,0.000028676302,0.0002724575,0.00051557575,0.0022459875,0.0026562077],"genre_scores_gemma":[0.9042853,0.00005960168,0.093575105,0.00003611861,0.000007517058,0.00016717744,0.0004739026,0.000060880997,0.0013343047],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99852645,0.00076648244,0.00007552612,0.00028883148,0.00024806277,0.00009465873],"domain_scores_gemma":[0.99575436,0.0022285508,0.00037872954,0.0003761813,0.0010169912,0.00024507593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032706156,0.0005071089,0.00035782615,0.0012972621,0.0004701676,0.0013255747,0.00080922595,0.00059713353,0.0013450306],"category_scores_gemma":[0.011108202,0.0002143234,0.0004119225,0.0006626207,0.00019763694,0.0015326126,0.0014519771,0.00093007454,0.00063587143],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056183623,0.00070179516,0.16953756,0.00017468164,0.00014416604,0.00023100499,0.006419628,0.1257361,0.012184393,0.0031207367,0.004286072,0.67690206],"study_design_scores_gemma":[0.000021570948,0.0003050473,0.06267007,0.000051903084,0.00007573607,0.00009236503,0.00259151,0.91072965,0.014104101,0.005886258,0.0033820188,0.00008977384],"about_ca_topic_score_codex":0.010833553,"about_ca_topic_score_gemma":0.020154376,"teacher_disagreement_score":0.010833553,"about_ca_system_score_codex":0.00092720566,"about_ca_system_score_gemma":0.0012693984,"threshold_uncertainty_score":0.021541},"labels":[],"label_agreement":null},{"id":"W4413082047","doi":"10.2139/ssrn.5339141","title":"Large Language Models Improve Hypothesis Generation by Reducing Effort","year":2025,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Quest University Canada; University of Toronto","funders":"","keywords":"Computer science; Linguistics; Philosophy","score_opus":0.014253216304133851,"score_gpt":0.24683500276385686,"score_spread":0.23258178645972302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413082047","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09741589,0.0031799695,0.8720504,0.00484171,0.0009869507,0.00086020917,0.0019517412,0.012793786,0.0059192996],"genre_scores_gemma":[0.63188016,0.0012421027,0.34008443,0.0032599734,0.0021350658,0.0011095081,0.008834507,0.0022943066,0.009159953],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9835384,0.011438533,0.0008151602,0.0022939998,0.0015380079,0.0003757519],"domain_scores_gemma":[0.75084794,0.2295779,0.0023278766,0.011746015,0.0039274343,0.0015728767],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.021104889,0.0034157515,0.0029012896,0.004961308,0.0016882281,0.004719095,0.0031898334,0.005145378,0.012545416],"category_scores_gemma":[0.13296406,0.0017238894,0.003222737,0.0029948498,0.0012177333,0.009168262,0.0049669724,0.0061878213,0.00906282],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040201494,0.0018867576,0.023367869,0.0010101816,0.0019583488,0.000855222,0.0006800335,0.10434086,0.018720614,0.0078440225,0.03455884,0.8007571],"study_design_scores_gemma":[0.0007511502,0.0006205958,0.004390478,0.00011165948,0.001324672,0.00034870234,0.00019978902,0.92646694,0.008642817,0.05133743,0.0057128854,0.00009293462],"about_ca_topic_score_codex":0.0016331823,"about_ca_topic_score_gemma":0.0038098365,"teacher_disagreement_score":0.9788951,"about_ca_system_score_codex":0.00079012976,"about_ca_system_score_gemma":0.0024714977,"threshold_uncertainty_score":0.111614645},"labels":[],"label_agreement":null},{"id":"W4413089949","doi":"10.1017/9781009636124.008","title":"Detecting Latent Signs","year":2025,"lang":"en","type":"book-chapter","venue":"Cambridge University Press eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Probabilistic latent semantic analysis; Computer science; Artificial intelligence","score_opus":0.030735349222091465,"score_gpt":0.1980880582019128,"score_spread":0.16735270897982135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413089949","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020173231,0.0016083919,0.9476225,0.0005708461,0.00024699504,0.00011635783,0.0015758417,0.007175979,0.020909825],"genre_scores_gemma":[0.3106925,0.0025368189,0.5679055,0.00037479805,0.00042979603,0.00028495383,0.013246348,0.0022227704,0.10230647],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99901235,0.00021688131,0.000041602387,0.00031442702,0.00031018202,0.00010451181],"domain_scores_gemma":[0.9983699,0.0006968963,0.00011255672,0.00043888693,0.00029071522,0.00009104809],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008352188,0.001072949,0.0010820064,0.0025385297,0.0007304964,0.0032573955,0.0011320367,0.0016714552,0.017264333],"category_scores_gemma":[0.0044132927,0.00084997795,0.0010132756,0.0019660746,0.00068958034,0.0033872908,0.00223555,0.0017465808,0.016905742],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022676741,0.00011928305,0.0025945795,0.00021568002,0.00007160244,0.00021472486,0.00025326796,0.006383971,0.036298733,0.02865222,0.040528663,0.8844404],"study_design_scores_gemma":[0.000069363756,0.00027104586,0.009278914,0.0002262208,0.00022103962,0.0016874147,0.00061970693,0.6457059,0.064597435,0.17149815,0.1057062,0.00011865976],"about_ca_topic_score_codex":0.0017772679,"about_ca_topic_score_gemma":0.0032400168,"teacher_disagreement_score":0.017264333,"about_ca_system_score_codex":0.00043604727,"about_ca_system_score_gemma":0.0006738154,"threshold_uncertainty_score":0.057754934},"labels":[],"label_agreement":null},{"id":"W4413115901","doi":"10.1098/rstb.2023.0499","title":"Re-evaluating Theory of Mind evaluation in large language models","year":2025,"lang":"en","type":"article","venue":"Philosophical Transactions of the Royal Society B Biological Sciences","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Theory of mind; Computer science; Cognitive science; Linguistics; Natural language processing; Psychology; Philosophy; Neuroscience; Cognition","score_opus":0.12278221242681205,"score_gpt":0.3560545157109268,"score_spread":0.23327230328411475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413115901","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29772067,0.0049981503,0.6337642,0.027708594,0.0005905599,0.00042344542,0.00070263,0.001965049,0.032126717],"genre_scores_gemma":[0.8749234,0.0004721063,0.121701755,0.0007994621,0.00014573803,0.00017497738,0.00044623337,0.00052418525,0.00081217877],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.92605793,0.05875145,0.0024536205,0.002502767,0.009290552,0.0009435776],"domain_scores_gemma":[0.6688031,0.2787719,0.008419112,0.020780507,0.019900242,0.0033250873],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07438702,0.0014353443,0.0024038272,0.0032592088,0.0011719804,0.010237507,0.0034109289,0.0029416364,0.0038963733],"category_scores_gemma":[0.35956743,0.0009193723,0.0014750318,0.0018195013,0.005422107,0.014790715,0.00888461,0.0043929457,0.00058301055],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015668478,0.00042271637,0.026367692,0.0021582388,0.0013027673,0.0004970301,0.011187274,0.22463419,0.0039631175,0.4937969,0.009965129,0.22413807],"study_design_scores_gemma":[0.00011431123,0.0002383637,0.0028538702,0.00041092437,0.0001535649,0.0001020589,0.0014081275,0.45798254,0.0022485165,0.52816653,0.0061887368,0.00013240926],"about_ca_topic_score_codex":0.0059266402,"about_ca_topic_score_gemma":0.0076468075,"teacher_disagreement_score":0.925613,"about_ca_system_score_codex":0.0063530873,"about_ca_system_score_gemma":0.0041343705,"threshold_uncertainty_score":0.39340085},"labels":[],"label_agreement":null},{"id":"W4413140769","doi":"10.1148/rg.240103","title":"Deep Learning Models Connecting Images and Text: A Primer for Radiologists","year":2025,"lang":"en","type":"article","venue":"Radiographics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure; Centre Hospitalier de l’Université de Montréal; Western University","funders":"Canadian Institutes of Health Research; Fonds de Recherche du Québec - Santé; Fondation de l'Association des radiologistes du Québec; Radiological Society of North America","keywords":"Medicine; Primer (cosmetics); Deep learning; Artificial intelligence; Medical physics; Radiology; Natural language processing; Computer vision","score_opus":0.025815869188352252,"score_gpt":0.2688098754725671,"score_spread":0.24299400628421483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413140769","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00080592756,0.02775224,0.9350075,0.027185848,0.0011107745,0.00007333856,0.0007844344,0.0028485535,0.0044314144],"genre_scores_gemma":[0.038288128,0.059960023,0.86506075,0.007401447,0.006187834,0.0008622976,0.0026832416,0.0015823926,0.017973969],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992804,0.00027062875,0.00008070429,0.00014843268,0.00018056977,0.00003938453],"domain_scores_gemma":[0.9951757,0.0029716874,0.0002280751,0.00043998027,0.0009077878,0.00027667102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037345437,0.0014228603,0.0008202404,0.0020921582,0.00039996774,0.0034885628,0.0034832773,0.0038903018,0.008165705],"category_scores_gemma":[0.011131115,0.0019535637,0.0011966713,0.0016051955,0.0018638702,0.008031566,0.0023263951,0.008345078,0.0071241595],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011303537,0.00016147805,0.0013496567,0.00097206223,0.00014664598,0.0002792779,0.00031330757,0.031000786,0.0028962442,0.112543754,0.2042871,0.64593667],"study_design_scores_gemma":[0.000030681975,0.00009860063,0.00073069276,0.00104938,0.00006343536,0.0007635972,0.00013443899,0.22050472,0.0037428366,0.42159817,0.35115075,0.00013274392],"about_ca_topic_score_codex":0.0033045043,"about_ca_topic_score_gemma":0.0040226905,"teacher_disagreement_score":0.008165705,"about_ca_system_score_codex":0.0013718499,"about_ca_system_score_gemma":0.0011987488,"threshold_uncertainty_score":0.027317047},"labels":[],"label_agreement":null},{"id":"W4413145065","doi":"10.1109/cvpr52734.2025.02419","title":"Advancing Generalizable Tumor Segmentation with Anomaly-Aware Open-Vocabulary Attention Maps and Frozen Foundation Diffusion Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Foundation (evidence); Computer science; Segmentation; Anomaly (physics); Artificial intelligence; Vocabulary; Diffusion; Natural language processing; Geography; Linguistics; Physics","score_opus":0.013445897939900099,"score_gpt":0.2531931918543566,"score_spread":0.23974729391445654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413145065","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023611356,0.00039380838,0.97223026,0.0003078734,0.000045023742,0.00005373006,0.00018046271,0.0023845662,0.00079295767],"genre_scores_gemma":[0.57456005,0.000687955,0.41312268,0.0007785485,0.00023469077,0.00019861772,0.0018231749,0.0012059809,0.007388423],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995753,0.00008003749,0.000017766384,0.00017811674,0.000083737825,0.00006495458],"domain_scores_gemma":[0.9988053,0.00062226737,0.00010976969,0.00020313691,0.00017267847,0.00008687407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011357431,0.0012631429,0.0010494706,0.0011001208,0.00042858234,0.0012981348,0.0023334813,0.0018488317,0.001836082],"category_scores_gemma":[0.0037024391,0.00064824655,0.0016223775,0.000761085,0.0011519958,0.0019718064,0.0022037863,0.0024252879,0.0008977909],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025026267,0.00013539831,0.0020441776,0.0002851308,0.00015424428,0.00031321406,0.00060384854,0.57457924,0.03715507,0.015535682,0.0071219476,0.36182177],"study_design_scores_gemma":[0.00000786087,0.00002672906,0.0001547419,0.00000836514,0.000011224806,0.00005782828,0.000023607758,0.9857538,0.0034321907,0.009835492,0.0006788752,0.000009309516],"about_ca_topic_score_codex":0.008902923,"about_ca_topic_score_gemma":0.011773227,"teacher_disagreement_score":0.008902923,"about_ca_system_score_codex":0.0011737915,"about_ca_system_score_gemma":0.0011075891,"threshold_uncertainty_score":0.017702222},"labels":[],"label_agreement":null},{"id":"W4413157588","doi":"10.1109/cvpr52734.2025.01239","title":"StoryGPT-V: Large Language Models as Consistent Story Visualizers","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science","score_opus":0.01910971149770561,"score_gpt":0.29659802929818313,"score_spread":0.2774883178004775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413157588","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028968638,0.00058845116,0.9450394,0.00052526593,0.000093228846,0.00016968821,0.0019576626,0.018870829,0.0037869278],"genre_scores_gemma":[0.40086338,0.0004927388,0.5796304,0.0004492353,0.00006234433,0.0007959586,0.007232552,0.0034393906,0.007033957],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99964,0.00014599957,0.00001600406,0.00011017389,0.000062346924,0.000025487592],"domain_scores_gemma":[0.99918514,0.0005398852,0.000044128188,0.000111686866,0.00007063274,0.000048603546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006279632,0.0010945617,0.00048681654,0.0005616024,0.00029729333,0.0013365861,0.00212466,0.001294132,0.00727136],"category_scores_gemma":[0.0039616376,0.0006236456,0.0012378248,0.00040592247,0.00052292366,0.00186624,0.0015451455,0.0016700125,0.0019057328],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003505015,0.00015437164,0.0016325878,0.00046200698,0.0001502848,0.00047539064,0.00073770655,0.77097076,0.013906597,0.03590481,0.019433424,0.15582162],"study_design_scores_gemma":[0.00002402546,0.000015426587,0.000045155186,0.000009453133,0.0000068414333,0.000028382357,0.00002066941,0.9882762,0.0013422625,0.0076961876,0.0025289648,0.000006425547],"about_ca_topic_score_codex":0.003917247,"about_ca_topic_score_gemma":0.0079610385,"teacher_disagreement_score":0.00727136,"about_ca_system_score_codex":0.0008627693,"about_ca_system_score_gemma":0.0005898482,"threshold_uncertainty_score":0.024325073},"labels":[],"label_agreement":null},{"id":"W4413243547","doi":"10.1145/3759453","title":"A Survey of Conversational Search","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Information Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Semantic search; Context (archaeology); Unison; Search engine; Natural language; Data science; Human–computer interaction; World Wide Web; Information retrieval; Artificial intelligence","score_opus":0.0433947417254156,"score_gpt":0.28226504420324516,"score_spread":0.23887030247782956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413243547","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0206901,0.7653645,0.08779073,0.0059234677,0.0009176594,0.00055451354,0.001759994,0.0013398129,0.11565909],"genre_scores_gemma":[0.18475333,0.67945313,0.086832926,0.0033062773,0.0022914642,0.00069946214,0.005853734,0.00062883284,0.036180828],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99790204,0.000663822,0.00023173254,0.00029810262,0.0006846209,0.00021970776],"domain_scores_gemma":[0.9942082,0.0039460347,0.0002546532,0.00033106055,0.0010506223,0.00020937195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021095362,0.000976689,0.0011194913,0.006117749,0.0014621873,0.0037851357,0.0018902041,0.0016361916,0.012535036],"category_scores_gemma":[0.009987826,0.0006061095,0.00083624007,0.008871275,0.0007391242,0.009207413,0.0020624294,0.0010933601,0.004308348],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020243379,0.00013488097,0.0042151874,0.006518687,0.0000886029,0.00022698205,0.0017491196,0.002355804,0.00131571,0.04448105,0.03982554,0.89888597],"study_design_scores_gemma":[0.000026369813,0.00017058568,0.005071667,0.0026122867,0.00016104663,0.002033196,0.0022946245,0.016022079,0.0017797607,0.02814687,0.9415833,0.00009812134],"about_ca_topic_score_codex":0.008032735,"about_ca_topic_score_gemma":0.005784577,"teacher_disagreement_score":0.012535036,"about_ca_system_score_codex":0.0019053274,"about_ca_system_score_gemma":0.0033606181,"threshold_uncertainty_score":0.041933894},"labels":[],"label_agreement":null},{"id":"W4413277286","doi":"10.1109/access.2025.3599832","title":"A Comprehensive Survey on LLM-Powered Recommender Systems: From Discriminative, Generative to Multi-Modal Paradigms","year":2025,"lang":"en","type":"article","venue":"IEEE Access","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Toronto Metropolitan University","keywords":"Computer science; Discriminative model; Recommender system; Modal; Generative grammar; Artificial intelligence; Machine learning; Information retrieval","score_opus":0.15360534976935264,"score_gpt":0.3817094154353873,"score_spread":0.22810406566603464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413277286","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018493239,0.36920443,0.5856191,0.0041938643,0.00062450237,0.00040053125,0.0012253193,0.0027141639,0.017524857],"genre_scores_gemma":[0.19461094,0.3490579,0.43708223,0.0026621937,0.00214248,0.0005599065,0.0030170002,0.0005311103,0.010336207],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980987,0.00070048444,0.00020714104,0.0003598801,0.0005422679,0.00009148835],"domain_scores_gemma":[0.9939016,0.0042080944,0.00014883789,0.0007956229,0.0008390543,0.0001068266],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00304517,0.001247738,0.0015874828,0.002547498,0.0007725456,0.0022866845,0.0020933263,0.001497257,0.0050035994],"category_scores_gemma":[0.011076234,0.0009740721,0.0012065058,0.005644729,0.00048718587,0.004277855,0.0015839753,0.0016558598,0.0028070349],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014676205,0.00016797107,0.00430571,0.004375887,0.00027577626,0.00008781306,0.0003879763,0.016019698,0.0024235619,0.028949734,0.021995056,0.92086405],"study_design_scores_gemma":[0.0001042575,0.0008995987,0.009890835,0.0037385875,0.00089458755,0.0020999024,0.0008820859,0.4468189,0.007885227,0.09073218,0.43566164,0.00039226512],"about_ca_topic_score_codex":0.0047266884,"about_ca_topic_score_gemma":0.0069778142,"teacher_disagreement_score":0.0050035994,"about_ca_system_score_codex":0.0009945953,"about_ca_system_score_gemma":0.0011378678,"threshold_uncertainty_score":0.016738713},"labels":[],"label_agreement":null},{"id":"W4413362641","doi":"10.3389/fninf.2025.1609077","title":"Large language models can extract metadata for annotation of human neuroimaging publications","year":2025,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institute on Drug Abuse","keywords":"Annotation; Metadata; Computer science; Information retrieval; USable; Natural language processing; Neuroimaging; Artificial intelligence; Data science; World Wide Web; Psychology","score_opus":0.02792275338671165,"score_gpt":0.29132097665357126,"score_spread":0.26339822326685963,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413362641","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057763528,0.003961272,0.7764947,0.00512131,0.0011645058,0.0008332354,0.031158809,0.10339381,0.020108799],"genre_scores_gemma":[0.31838337,0.0014514821,0.600793,0.0014033045,0.000483715,0.0012426175,0.06250621,0.006424419,0.007311879],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98741376,0.0070812115,0.0011777554,0.0018419063,0.002140484,0.00034486112],"domain_scores_gemma":[0.9267244,0.049216688,0.0038542503,0.011867455,0.007429457,0.0009077902],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.018251741,0.002229901,0.00076090585,0.0081759915,0.001334242,0.005304558,0.0020802903,0.002082159,0.007430908],"category_scores_gemma":[0.09465841,0.00085990335,0.0017913698,0.0046455995,0.0010231847,0.00849527,0.004686649,0.0024345312,0.011916733],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015391612,0.0004641023,0.025682978,0.0052558924,0.0009582779,0.00083571166,0.004928236,0.034426313,0.027748995,0.028374555,0.19268064,0.6771051],"study_design_scores_gemma":[0.00023250918,0.00045501254,0.017348263,0.0012461822,0.0005529383,0.0009868714,0.0028278427,0.5419012,0.053934395,0.12224322,0.25777653,0.00049519714],"about_ca_topic_score_codex":0.004830987,"about_ca_topic_score_gemma":0.015899533,"teacher_disagreement_score":0.9817483,"about_ca_system_score_codex":0.0017363516,"about_ca_system_score_gemma":0.0039767264,"threshold_uncertainty_score":0.09652561},"labels":[],"label_agreement":null},{"id":"W4413364391","doi":"10.1016/j.procs.2025.07.178","title":"Retrieve-Classify-Read: Passage Filtering via Subject Classification for University Question Answering","year":2025,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Acadia University","funders":"","keywords":"Computer science; Question answering; Subject (documents); Information retrieval; Artificial intelligence; Natural language processing; World Wide Web","score_opus":0.022977001796318477,"score_gpt":0.2571357960003268,"score_spread":0.23415879420400834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413364391","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.088576674,0.0012477381,0.8479666,0.000586505,0.00028450665,0.0011786132,0.0028664144,0.054171503,0.0031215092],"genre_scores_gemma":[0.38116562,0.000526413,0.597157,0.0003814061,0.00025025703,0.00063231273,0.01096343,0.0011046798,0.007818881],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985708,0.00058673305,0.00010702621,0.00040898132,0.0002389161,0.00008757078],"domain_scores_gemma":[0.99537,0.0025089174,0.00023992792,0.0008935862,0.0007410605,0.0002465218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032796906,0.0012963401,0.0010903869,0.002764071,0.0007745314,0.0014307179,0.0020527802,0.0017466057,0.005609877],"category_scores_gemma":[0.011299276,0.00031559065,0.0015379209,0.0012378613,0.0005712264,0.0024062344,0.0014614626,0.001540485,0.0039401352],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010430069,0.00094529474,0.01276717,0.00057905517,0.00029615295,0.0003682858,0.0013298,0.029206881,0.03176753,0.0037375314,0.024351541,0.8936077],"study_design_scores_gemma":[0.00011576981,0.0008491046,0.008450267,0.00005355728,0.00019161409,0.00039996253,0.00036551344,0.92013425,0.03653907,0.0072544315,0.025546793,0.000099690515],"about_ca_topic_score_codex":0.013904609,"about_ca_topic_score_gemma":0.01265228,"teacher_disagreement_score":0.013904609,"about_ca_system_score_codex":0.0011292681,"about_ca_system_score_gemma":0.0015982304,"threshold_uncertainty_score":0.027647316},"labels":[],"label_agreement":null},{"id":"W4413368058","doi":"10.1007/978-3-031-90573-5_3","title":"Architectural Deep Dive into Large Language Models","year":2025,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing","score_opus":0.06410466795799205,"score_gpt":0.3517674512652404,"score_spread":0.28766278330724837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413368058","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0052026454,0.0010415531,0.9571454,0.0022572982,0.00016569908,0.000047727157,0.00054731173,0.005507024,0.02808526],"genre_scores_gemma":[0.17001696,0.0031269437,0.748627,0.0016945297,0.0004332354,0.00021637788,0.003561474,0.0058880053,0.06643549],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99942315,0.0001947881,0.000038652197,0.00011685909,0.00018665305,0.00003998037],"domain_scores_gemma":[0.9979316,0.0010879285,0.00004181085,0.0007211116,0.00014733174,0.000070307586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010265938,0.00083936914,0.00071818545,0.0007187692,0.00072486757,0.0030223844,0.0015757731,0.00088690204,0.01799924],"category_scores_gemma":[0.0047019124,0.0011287562,0.0014881939,0.0012667822,0.0014300717,0.010773473,0.002916403,0.0044131605,0.008802701],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006563213,0.00005604623,0.00044845903,0.00023011386,0.000055055054,0.0001250684,0.0006875207,0.01594184,0.004224558,0.7402226,0.037034,0.20090918],"study_design_scores_gemma":[0.000007650416,0.000017141749,0.0000973277,0.000048299116,0.000034476503,0.0001242109,0.00012499912,0.13612194,0.0025658552,0.76241136,0.098432675,0.0000138800115],"about_ca_topic_score_codex":0.0022095817,"about_ca_topic_score_gemma":0.005617263,"teacher_disagreement_score":0.01799924,"about_ca_system_score_codex":0.0009494572,"about_ca_system_score_gemma":0.0012177054,"threshold_uncertainty_score":0.060213447},"labels":[],"label_agreement":null},{"id":"W4413411773","doi":"10.1145/3721145.3730418","title":"Cephalo: Harnessing Heterogeneous GPU Clusters for Training Transformer Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Transformer; Training (meteorology); Supercomputer; Parallel computing; Electrical engineering; Engineering; Physics; Voltage","score_opus":0.05453842143186296,"score_gpt":0.2824619046238849,"score_spread":0.2279234831920219,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413411773","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06511579,0.0016655955,0.81858534,0.0006741225,0.00073798577,0.00030380065,0.0049230475,0.098453626,0.009540683],"genre_scores_gemma":[0.42988575,0.00079827284,0.5310238,0.00058860355,0.00017804216,0.00043130686,0.017743114,0.006741452,0.012609619],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99961406,0.00008836616,0.000022060025,0.00013846323,0.00008118599,0.000055981567],"domain_scores_gemma":[0.999421,0.00022962564,0.000022307664,0.00016363862,0.000108256034,0.000055236436],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00063530053,0.0014955428,0.0007468805,0.0009841912,0.0005691566,0.0014633427,0.0020398544,0.0009307057,0.010389338],"category_scores_gemma":[0.0029336235,0.0007626608,0.0010281183,0.0014324405,0.0002883677,0.0016520307,0.0013431516,0.0018849946,0.0052209897],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010699448,0.0003132103,0.0044943825,0.00037942897,0.00042470428,0.0003000491,0.00029243148,0.20834014,0.016296439,0.0063987924,0.11596301,0.6457276],"study_design_scores_gemma":[0.00008335399,0.00005254679,0.00033629965,0.000011842489,0.00003910872,0.000043394768,0.000059797574,0.9802741,0.0054958556,0.004448761,0.009142227,0.000012644928],"about_ca_topic_score_codex":0.0147958435,"about_ca_topic_score_gemma":0.037908588,"teacher_disagreement_score":0.0147958435,"about_ca_system_score_codex":0.0007118854,"about_ca_system_score_gemma":0.0014394857,"threshold_uncertainty_score":0.034755766},"labels":[],"label_agreement":null},{"id":"W4413423557","doi":"10.1016/b978-0-443-30046-2.00012-0","title":"Identifying large language model hallucinations in health communication","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Psychology; Computer science","score_opus":0.0335838217936583,"score_gpt":0.3080661655594898,"score_spread":0.27448234376583147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423557","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037905283,0.0053673107,0.9410103,0.0019971184,0.00036726618,0.00007386474,0.0015621732,0.0026531261,0.009063478],"genre_scores_gemma":[0.59468395,0.0081306435,0.342937,0.00082169,0.001133837,0.00028978553,0.007541334,0.0009178082,0.043543942],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99957806,0.00020035541,0.000022333543,0.000085397376,0.00008213087,0.00003177171],"domain_scores_gemma":[0.9971806,0.0024548739,0.0000784131,0.00013126928,0.00010336395,0.000051358606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012666875,0.0010142144,0.00069939485,0.00095505104,0.00034851505,0.0017427806,0.0006314793,0.0010916231,0.0064121787],"category_scores_gemma":[0.00471819,0.00038353502,0.00089545734,0.0011176468,0.0003563112,0.0015632154,0.001009802,0.0015205136,0.0036083763],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000276695,0.00018966256,0.0053713443,0.00041954225,0.00021023404,0.00048969535,0.0006607064,0.074445345,0.015737198,0.014740386,0.034293216,0.853166],"study_design_scores_gemma":[0.000016776157,0.000109817214,0.005205263,0.000107416716,0.00007343356,0.0005521369,0.00040681634,0.9293423,0.0055410503,0.045474257,0.013123473,0.000047236208],"about_ca_topic_score_codex":0.0027922355,"about_ca_topic_score_gemma":0.0032383741,"teacher_disagreement_score":0.0064121787,"about_ca_system_score_codex":0.00039321088,"about_ca_system_score_gemma":0.0004121521,"threshold_uncertainty_score":0.021450877},"labels":[],"label_agreement":null},{"id":"W4413423575","doi":"10.1016/b978-0-443-30046-2.00007-7","title":"Leveraging medical discourse to answer complex questions","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Linguistics; Sociology; Epistemology; Philosophy","score_opus":0.03216196258257568,"score_gpt":0.29611953986730616,"score_spread":0.26395757728473046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423575","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0128948465,0.008022674,0.88915205,0.010334861,0.00081546605,0.00018256917,0.0017422825,0.003670236,0.07318505],"genre_scores_gemma":[0.22704461,0.008914671,0.7036829,0.002174059,0.0016867945,0.00031927248,0.0055950605,0.0012611562,0.049321465],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984939,0.0008079874,0.00008951831,0.0002160697,0.00034593497,0.00004659365],"domain_scores_gemma":[0.99193656,0.007168138,0.00019000014,0.00030634904,0.00029233197,0.00010666614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028789693,0.0010237347,0.00046578373,0.0023917241,0.00059822644,0.005347528,0.0008933018,0.0016974008,0.01582503],"category_scores_gemma":[0.011990589,0.00045655872,0.00084352924,0.0016733484,0.0011349709,0.0060882727,0.0026392557,0.0017975946,0.0054158852],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010504626,0.00013225064,0.0016877288,0.0011951993,0.00010415506,0.00035516385,0.005722946,0.010093438,0.014008551,0.17352563,0.054668076,0.73840183],"study_design_scores_gemma":[0.000032686505,0.00008090866,0.0014875174,0.00081216625,0.00012975515,0.00049696275,0.00267325,0.14153366,0.008859879,0.45967183,0.38414964,0.00007169793],"about_ca_topic_score_codex":0.0012836477,"about_ca_topic_score_gemma":0.0016128722,"teacher_disagreement_score":0.01582503,"about_ca_system_score_codex":0.00074033026,"about_ca_system_score_gemma":0.00079735427,"threshold_uncertainty_score":0.05294001},"labels":[],"label_agreement":null},{"id":"W4413423589","doi":"10.1016/b978-0-443-30046-2.00010-7","title":"Differential diagnosis making with large language models and probabilistic logic program","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Probabilistic logic; Computer science; Differential (mechanical device); Programming language; Artificial intelligence; Engineering","score_opus":0.022528924438561714,"score_gpt":0.2635118232668709,"score_spread":0.24098289882830917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423589","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009690299,0.0007487786,0.9764291,0.0012636977,0.00007195891,0.000027602287,0.00018215895,0.0006994597,0.010886834],"genre_scores_gemma":[0.4518315,0.0016887903,0.51468796,0.00045835803,0.00039404366,0.00013814699,0.00083402183,0.00029630217,0.029670814],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99922025,0.0002210436,0.00004066264,0.00017359876,0.00030056882,0.00004397275],"domain_scores_gemma":[0.99436,0.0049803746,0.00013478768,0.00026271146,0.00020418219,0.000057929734],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010693265,0.0005923812,0.00059425074,0.0008091446,0.00050106464,0.0017384993,0.0011948335,0.0007122023,0.007262173],"category_scores_gemma":[0.0058575864,0.0005392668,0.0009329049,0.0008608986,0.0012290622,0.0038222247,0.0015365448,0.0020928592,0.00094864704],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000123941,0.00009965048,0.0012721829,0.0003258458,0.000066136796,0.00050456764,0.00038185765,0.16914572,0.0027879307,0.47655916,0.015038443,0.33369467],"study_design_scores_gemma":[0.000008117308,0.0000133531175,0.00017339159,0.000020745658,0.000021471604,0.00017324962,0.000041926207,0.46239045,0.0012028035,0.5292273,0.006715737,0.00001145376],"about_ca_topic_score_codex":0.0019268638,"about_ca_topic_score_gemma":0.0026216367,"teacher_disagreement_score":0.007262173,"about_ca_system_score_codex":0.0013289454,"about_ca_system_score_gemma":0.00083587924,"threshold_uncertainty_score":0.024294376},"labels":[],"label_agreement":null},{"id":"W4413423613","doi":"10.1016/b978-0-443-30046-2.00011-9","title":"Enabling large language models with explainability","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Linguistics; Philosophy","score_opus":0.01574175996363337,"score_gpt":0.2376405420623077,"score_spread":0.22189878209867434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423613","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022046864,0.0006150155,0.97468525,0.0009432388,0.00011714863,0.000030185407,0.00052493816,0.006250767,0.014628771],"genre_scores_gemma":[0.17861919,0.0041890643,0.7373932,0.0007305846,0.0005154084,0.00035243496,0.0060644043,0.007770321,0.06436541],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915767,0.000290013,0.000044544675,0.00017793215,0.000273177,0.00005657512],"domain_scores_gemma":[0.99712425,0.0021529756,0.00006106025,0.00048375176,0.0001238203,0.00005410588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001125978,0.0010842888,0.00068274944,0.0008168657,0.00051693263,0.0033074094,0.0013685041,0.001146332,0.032743677],"category_scores_gemma":[0.006203376,0.0011099037,0.0017256201,0.0011400833,0.00090170663,0.006969744,0.0034025381,0.0032954314,0.010616815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007938314,0.0000733689,0.0006004582,0.00049604336,0.00009535993,0.00027873932,0.00059074373,0.043585036,0.0069819973,0.55459774,0.0582005,0.33442065],"study_design_scores_gemma":[0.000017919636,0.000016735497,0.00017696964,0.00009045037,0.000052353,0.00018833649,0.00009755852,0.27773586,0.005280112,0.6002926,0.1160216,0.000029456005],"about_ca_topic_score_codex":0.0023387468,"about_ca_topic_score_gemma":0.0024853472,"teacher_disagreement_score":0.032743677,"about_ca_system_score_codex":0.0007923031,"about_ca_system_score_gemma":0.00070887175,"threshold_uncertainty_score":0.109538436},"labels":[],"label_agreement":null},{"id":"W4413423668","doi":"10.1016/b978-0-443-30046-2.00015-6","title":"Kolmogorov–Arnold network for word-level explainable meaning representation","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Meaning (existential); Word (group theory); Representation (politics); Linguistics; Natural language processing; Computer science; Psychology; Philosophy; Political science; Psychotherapist","score_opus":0.04956949208312871,"score_gpt":0.2724836734458199,"score_spread":0.2229141813626912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423668","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010848182,0.0013190305,0.97546357,0.0008997353,0.0001370577,0.000023298508,0.0004921936,0.0005764828,0.010240513],"genre_scores_gemma":[0.5875002,0.0039472394,0.3746521,0.00030259005,0.0005110229,0.00026189836,0.002392903,0.00052700826,0.029905004],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997458,0.00008632486,0.000023211805,0.000069605776,0.000055581797,0.000019483285],"domain_scores_gemma":[0.9991518,0.00057785196,0.000053736407,0.00011730833,0.0000694614,0.000029854415],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00047795134,0.0005126599,0.00053286453,0.001093558,0.00040237396,0.0015061693,0.0008904821,0.0010676429,0.008845976],"category_scores_gemma":[0.0029584793,0.0003419412,0.0008217953,0.0011946629,0.00083732174,0.0037906866,0.0011536562,0.0015823398,0.002051949],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055857603,0.000030319336,0.0006581531,0.0001821851,0.00006155193,0.00015780407,0.00021099628,0.079379074,0.0023909386,0.7707087,0.007939928,0.1382245],"study_design_scores_gemma":[0.0000030502476,0.000006497787,0.00016590997,0.000018749864,0.000010157432,0.000049160102,0.000020702339,0.34242117,0.00040264177,0.65228134,0.0046098437,0.000010818647],"about_ca_topic_score_codex":0.0016478389,"about_ca_topic_score_gemma":0.0015759944,"teacher_disagreement_score":0.008845976,"about_ca_system_score_codex":0.0008293588,"about_ca_system_score_gemma":0.0004950539,"threshold_uncertainty_score":0.029592752},"labels":[],"label_agreement":null},{"id":"W4413423671","doi":"10.1016/b978-0-443-30046-2.00014-4","title":"Enabling large language model with plug-and-play symbolic reasoning components","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Programming language; Symbolic execution; Cognitive science; Natural language processing; Psychology; Software","score_opus":0.013852066635577465,"score_gpt":0.2351449695671962,"score_spread":0.2212929029316187,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423671","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025669136,0.00024266027,0.95823896,0.00025298193,0.00009256195,0.00007902881,0.000406801,0.02104852,0.017071486],"genre_scores_gemma":[0.1207792,0.0015021213,0.81486535,0.00033738025,0.00011122834,0.00037998412,0.0037642387,0.011384933,0.0468756],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99894494,0.0002066077,0.000072470255,0.00018406991,0.00048099522,0.00011093518],"domain_scores_gemma":[0.99862754,0.0007869336,0.00004105564,0.00035088215,0.0001303241,0.00006321039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010785306,0.0016294585,0.0008733664,0.0010242478,0.00086242164,0.0061609168,0.002950059,0.0015910737,0.036553048],"category_scores_gemma":[0.0037634287,0.0014700888,0.0019523212,0.0011365343,0.0016311201,0.007401756,0.0043288316,0.0037335732,0.012249563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003324903,0.00023496644,0.00067158166,0.0010073475,0.000108386645,0.00060897676,0.0016770842,0.052592862,0.027613342,0.5355228,0.050043594,0.32958654],"study_design_scores_gemma":[0.00009110094,0.000056964458,0.00014473652,0.0002445284,0.00012819046,0.00042957533,0.00030311313,0.40414628,0.034756783,0.3083907,0.25123537,0.0000725717],"about_ca_topic_score_codex":0.0048682303,"about_ca_topic_score_gemma":0.0069369073,"teacher_disagreement_score":0.036553048,"about_ca_system_score_codex":0.0014175612,"about_ca_system_score_gemma":0.0017529767,"threshold_uncertainty_score":0.12228215},"labels":[],"label_agreement":null},{"id":"W4413423694","doi":"10.1016/b978-0-443-30046-2.00001-6","title":"Explainability discourse","year":2025,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Sociology; Linguistics; History; Philosophy","score_opus":0.017836107161482845,"score_gpt":0.26214607929721173,"score_spread":0.2443099721357289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413423694","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005514953,0.00818791,0.037427377,0.009284113,0.00058659154,0.000028427392,0.00042200633,0.00039040597,0.9381583],"genre_scores_gemma":[0.33253884,0.010077215,0.014112642,0.0010806666,0.001751498,0.00013427503,0.0021273235,0.0010812258,0.63709617],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99943274,0.0002516122,0.000020490008,0.00010826701,0.0001499,0.00003696261],"domain_scores_gemma":[0.99893004,0.0007684808,0.000044652115,0.00013110002,0.00009336712,0.000032306903],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007508103,0.00080277177,0.00038844475,0.0018833175,0.0013830173,0.005136233,0.00057210884,0.0013304441,0.05844756],"category_scores_gemma":[0.0031299146,0.00031567126,0.0003757802,0.0018616988,0.0025859172,0.006398071,0.0018656638,0.0019129034,0.008019763],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010677782,0.000010323149,0.00013884978,0.00011421307,0.0000053698577,0.000045501325,0.003505425,0.00032407066,0.00038159647,0.8843954,0.045533326,0.06553514],"study_design_scores_gemma":[0.0000048508377,0.0000061997675,0.00039846977,0.00020625166,0.000006911646,0.0000868156,0.0012695593,0.0010911421,0.00039532207,0.3868464,0.6096812,0.0000068998306],"about_ca_topic_score_codex":0.0027307097,"about_ca_topic_score_gemma":0.0025446923,"teacher_disagreement_score":0.05844756,"about_ca_system_score_codex":0.0023847646,"about_ca_system_score_gemma":0.00085552846,"threshold_uncertainty_score":0.19552654},"labels":[],"label_agreement":null},{"id":"W4413426951","doi":"10.21203/rs.3.rs-6959723/v1","title":"Retrieval-Augmented Generation for Natural Language Processing: A Survey","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Natural (archaeology); Computer science; Natural language processing; Information retrieval; Geography","score_opus":0.16078360183959114,"score_gpt":0.4509311062684964,"score_spread":0.29014750442890525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413426951","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012066018,0.12137735,0.84824,0.0013527947,0.0004776041,0.00027106528,0.0011278731,0.008126586,0.0069606192],"genre_scores_gemma":[0.19712092,0.10883374,0.66795397,0.0012834263,0.0019356986,0.00058113114,0.007762857,0.0021258679,0.012402404],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.997603,0.00084088056,0.0001906023,0.0007038505,0.0005466032,0.000115058174],"domain_scores_gemma":[0.9945262,0.0037117898,0.00012048854,0.00095697696,0.0006001152,0.00008437382],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028646868,0.0012051786,0.002380642,0.0028435518,0.0005833096,0.0025089453,0.0034532903,0.0015230573,0.006974498],"category_scores_gemma":[0.008399905,0.0008260171,0.0016933542,0.0037663728,0.0007627168,0.0046622385,0.0018308955,0.0014740763,0.0040413984],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014254489,0.00022320461,0.0006813965,0.0015742201,0.000087346016,0.00005072029,0.00015233127,0.009759792,0.003503313,0.009733946,0.0126191825,0.9614719],"study_design_scores_gemma":[0.00014038148,0.00067280524,0.0036996969,0.00072997605,0.00047179384,0.0013004235,0.00053773803,0.66123086,0.027129237,0.108994365,0.19491033,0.00018245958],"about_ca_topic_score_codex":0.0044933753,"about_ca_topic_score_gemma":0.0043116244,"teacher_disagreement_score":0.006974498,"about_ca_system_score_codex":0.00080856,"about_ca_system_score_gemma":0.0017837876,"threshold_uncertainty_score":0.023332},"labels":[],"label_agreement":null},{"id":"W4413451213","doi":"10.1016/j.jss.2025.112594","title":"Hybrid approach for multilevel multi-class requirement classification: Impact of stop-word removal and data augmentation","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Natural Sciences and Engineering Research Council of Canada; University of Saskatchewan","keywords":"Class (philosophy); Word (group theory); Computer science; Artificial intelligence; Data mining; Natural language processing; Engineering; Mathematics","score_opus":0.13341908526109603,"score_gpt":0.36219348799052786,"score_spread":0.22877440272943184,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413451213","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13702412,0.0014324946,0.845178,0.000918382,0.0004069218,0.00039911878,0.0021384275,0.009457503,0.003045044],"genre_scores_gemma":[0.44290417,0.00032186386,0.5442976,0.00042446767,0.00020324426,0.0004261532,0.006080738,0.0005832149,0.0047585773],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976307,0.00057638105,0.00022704482,0.00065589446,0.0006200555,0.00028985823],"domain_scores_gemma":[0.9942391,0.0028898802,0.00021377743,0.00082849606,0.0015883083,0.00024043207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002592547,0.0011695409,0.0018118204,0.002532689,0.00092447014,0.0020567256,0.0026191655,0.0018318796,0.0039121895],"category_scores_gemma":[0.00646647,0.0004064415,0.0020242387,0.0024246282,0.00034388978,0.0023546864,0.0017935972,0.0026686531,0.0030653966],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014622447,0.0017386052,0.010388946,0.00052820536,0.0005098797,0.0002008878,0.00054715807,0.02875733,0.034726217,0.0020864822,0.012088315,0.9069658],"study_design_scores_gemma":[0.00007082018,0.00022293668,0.005211799,0.00004194744,0.0002213005,0.00014000869,0.00040467447,0.9703718,0.01472739,0.0029942028,0.0055346927,0.00005829249],"about_ca_topic_score_codex":0.009497127,"about_ca_topic_score_gemma":0.018932737,"teacher_disagreement_score":0.009497127,"about_ca_system_score_codex":0.00056725513,"about_ca_system_score_gemma":0.0023859944,"threshold_uncertainty_score":0.018883705},"labels":[],"label_agreement":null},{"id":"W4413458163","doi":"10.1109/icdew67478.2025.00018","title":"LLM + Vector Data: Coupling of Large Language Models with Vector Data Management for Enhancing Data Science","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Vector (molecular biology); Coupling (piping); Data modeling; Data mining; Database; Engineering; Biology","score_opus":0.06124829965169441,"score_gpt":0.33251589297921424,"score_spread":0.2712675933275198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413458163","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004601193,0.0006947286,0.983493,0.002361952,0.00025785272,0.00010019536,0.0006881796,0.0061447173,0.0016581317],"genre_scores_gemma":[0.14721355,0.0014121906,0.83608484,0.0016884567,0.00056330726,0.00054764614,0.004496216,0.001986388,0.006007408],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99545467,0.002269452,0.00037117774,0.000849226,0.0008550585,0.00020038857],"domain_scores_gemma":[0.9906508,0.004465155,0.00038491117,0.0031687564,0.0009097646,0.00042065303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006606678,0.0011174887,0.0011288744,0.0013851158,0.0007491572,0.004959524,0.0027902208,0.0014387395,0.0063692206],"category_scores_gemma":[0.021312661,0.00091741444,0.0018734549,0.0020624273,0.0016670974,0.011205081,0.0067419587,0.0046860417,0.003898732],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043458497,0.00043385406,0.00386819,0.0005056764,0.0002915696,0.00027066935,0.0013522687,0.1159946,0.009093191,0.25857386,0.07169275,0.5374888],"study_design_scores_gemma":[0.00003745835,0.0000755064,0.00030768485,0.00005675959,0.000030604384,0.00006641611,0.00013229449,0.7619957,0.0041865995,0.19306849,0.039986674,0.00005583167],"about_ca_topic_score_codex":0.0050056893,"about_ca_topic_score_gemma":0.0075964956,"teacher_disagreement_score":0.006606678,"about_ca_system_score_codex":0.0018102016,"about_ca_system_score_gemma":0.002567028,"threshold_uncertainty_score":0.034939885},"labels":[],"label_agreement":null},{"id":"W4413467640","doi":"10.2196/76252","title":"Automated Literature Screening for Hepatocellular Carcinoma Treatment Through Integration of 3 Large Language Models: Methodological Study","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Hepatocellular carcinoma; Computer science; Medicine; World Wide Web; Internal medicine","score_opus":0.10811430115678947,"score_gpt":0.3935581168012115,"score_spread":0.285443815644422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413467640","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054701496,0.012396371,0.8889663,0.007505443,0.00032814004,0.0020485413,0.01530649,0.015294006,0.003453164],"genre_scores_gemma":[0.19545713,0.0014326569,0.7915737,0.0012760579,0.00010943599,0.0014263252,0.0077268207,0.00037296998,0.0006248863],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9812842,0.012957162,0.0021714866,0.0019201204,0.001464957,0.00020212731],"domain_scores_gemma":[0.91325414,0.07208654,0.0043856828,0.004474132,0.0050215395,0.0007780379],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.033887032,0.001588762,0.001735094,0.0067280415,0.0011038155,0.003934794,0.002360961,0.0016737801,0.004988311],"category_scores_gemma":[0.13416837,0.0010060171,0.005364451,0.0046280464,0.0006044397,0.0025823798,0.0040404927,0.0017051486,0.001226923],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024338714,0.00051210943,0.058175556,0.010731093,0.0066769794,0.0008861125,0.0014623073,0.1462736,0.004827707,0.012985749,0.03233455,0.7227004],"study_design_scores_gemma":[0.00094425515,0.0003230418,0.007165714,0.0016146945,0.0031879568,0.0005527302,0.00043417324,0.90636444,0.0051588346,0.048561227,0.025464723,0.00022811347],"about_ca_topic_score_codex":0.013641698,"about_ca_topic_score_gemma":0.028335065,"teacher_disagreement_score":0.966113,"about_ca_system_score_codex":0.0024783446,"about_ca_system_score_gemma":0.011024629,"threshold_uncertainty_score":0.17921388},"labels":[],"label_agreement":null},{"id":"W4413676882","doi":"10.64628/aam.4k53chh7v","title":"Why AI can’t take over creative writing","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Creative writing; Psychology; Computer science; Art; Literature","score_opus":0.02924258142513213,"score_gpt":0.29289440024972335,"score_spread":0.2636518188245912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413676882","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035129014,0.0104903,0.13168325,0.3395145,0.009277332,0.00014796942,0.00054398144,0.003454706,0.46975893],"genre_scores_gemma":[0.79111665,0.002830464,0.038634382,0.026408477,0.0032934768,0.00021714943,0.00051173224,0.0023319595,0.13465565],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98890126,0.005535345,0.00047625296,0.0013486443,0.0028416729,0.00089682976],"domain_scores_gemma":[0.9412334,0.037054833,0.00180492,0.0080377525,0.008469242,0.0033999048],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011929143,0.00056186685,0.0005363766,0.0016392752,0.0035637415,0.012091582,0.0014251276,0.004118342,0.024222882],"category_scores_gemma":[0.07758629,0.00055106817,0.00075238,0.0020041047,0.006167074,0.017717931,0.003684774,0.0065766214,0.01215196],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023516569,0.00012680383,0.0038005142,0.00044532013,0.00009400525,0.00026882312,0.011700754,0.0010487003,0.001661868,0.6438492,0.15250832,0.18426056],"study_design_scores_gemma":[0.000064917724,0.000026450281,0.0010037662,0.00019669015,0.000027892473,0.00032901493,0.005134808,0.0043584094,0.0013219379,0.7301654,0.2573287,0.000041844258],"about_ca_topic_score_codex":0.0038365526,"about_ca_topic_score_gemma":0.0021023594,"teacher_disagreement_score":0.024222882,"about_ca_system_score_codex":0.0019584806,"about_ca_system_score_gemma":0.0018843594,"threshold_uncertainty_score":0.08103359},"labels":[],"label_agreement":null},{"id":"W4413726998","doi":"10.1007/978-3-031-99854-6_10","title":"Leveraging Expert Usage to Speed up LLM Inference with Expert Parallelism","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"International Laboratory for Brain, Music and Sound Research","funders":"","keywords":"Computer science; Parallelism (grammar); Inference; Speedup; Parallel computing; Data parallelism; Expert system; Artificial intelligence","score_opus":0.03508971236646834,"score_gpt":0.2791464190751583,"score_spread":0.24405670670868995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413726998","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026786976,0.00042090917,0.9483948,0.0005822548,0.00020402364,0.00012327191,0.00046141387,0.014498832,0.008527627],"genre_scores_gemma":[0.3489401,0.00022510775,0.6365342,0.0006468982,0.0002571726,0.00013057892,0.0014424019,0.0017122022,0.0101113375],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99802434,0.00061197556,0.00012678638,0.000558696,0.00048584203,0.00019233623],"domain_scores_gemma":[0.9931492,0.0038986637,0.00019089613,0.0018585112,0.00067449233,0.00022833269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021015322,0.0010072979,0.0010775862,0.0011764388,0.000676445,0.0018773396,0.002556289,0.0015378096,0.0187921],"category_scores_gemma":[0.014119913,0.0008386057,0.0010823034,0.0011862823,0.00066902203,0.004909842,0.0029376668,0.0026422485,0.0062983776],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008463789,0.00041045708,0.004161647,0.0002618783,0.00015913349,0.00022511792,0.00032057374,0.06554599,0.018544523,0.016295109,0.027294843,0.86593443],"study_design_scores_gemma":[0.00007890717,0.000059348047,0.00065193453,0.000024249957,0.000048429873,0.00007272813,0.000060664202,0.94815683,0.00768377,0.0362121,0.0069323457,0.000018652265],"about_ca_topic_score_codex":0.005319994,"about_ca_topic_score_gemma":0.015646145,"teacher_disagreement_score":0.0187921,"about_ca_system_score_codex":0.0007853871,"about_ca_system_score_gemma":0.0016890148,"threshold_uncertainty_score":0.06286579},"labels":[],"label_agreement":null},{"id":"W4413741796","doi":"10.1111/2041-210x.70120","title":"New frontiers in artificial intelligence for biodiversity research and conservation with multimodal language models","year":2025,"lang":"en","type":"article","venue":"Methods in Ecology and Evolution","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Directorate for Biological Sciences; National Science Foundation","keywords":"Biodiversity; Biodiversity conservation; Computer science; Data science; Artificial intelligence; Ecology; Geography; Biology","score_opus":0.10758836602062251,"score_gpt":0.4065905817476293,"score_spread":0.2990022157270068,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413741796","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015693864,0.013278198,0.8664335,0.07386055,0.00085623335,0.00011349541,0.00047682843,0.0007202263,0.028567048],"genre_scores_gemma":[0.46846136,0.014055652,0.4972521,0.009689252,0.0018531568,0.00053633144,0.000868862,0.00031452838,0.0069687907],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9958066,0.0031249612,0.00013983597,0.00031612997,0.00047611108,0.00013631258],"domain_scores_gemma":[0.9810276,0.016030489,0.0003709052,0.0010368036,0.0010886653,0.00044547606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008797431,0.00085202086,0.0008227947,0.0014414082,0.0007425726,0.0058196564,0.0018951936,0.0028203952,0.0063832137],"category_scores_gemma":[0.017651532,0.0004803414,0.002105489,0.001131393,0.0049085002,0.010239532,0.0039335852,0.0056394106,0.0010929445],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008685253,0.00013356644,0.0019412511,0.00055473,0.00015771811,0.00030249057,0.0011998366,0.0646522,0.0015759872,0.8138184,0.016259132,0.0993178],"study_design_scores_gemma":[0.000013768195,0.000030595376,0.0003358827,0.00028509882,0.000021928867,0.00008445866,0.00029500443,0.20545949,0.000356274,0.7673457,0.025723666,0.00004820412],"about_ca_topic_score_codex":0.0038669987,"about_ca_topic_score_gemma":0.0033639004,"teacher_disagreement_score":0.008797431,"about_ca_system_score_codex":0.0019465967,"about_ca_system_score_gemma":0.0016838891,"threshold_uncertainty_score":0.046525836},"labels":[],"label_agreement":null},{"id":"W4413794792","doi":"10.3390/make7030089","title":"AlzheimerRAG: Multimodal Retrieval-Augmented Generation for Clinical Use Cases","year":2025,"lang":"en","type":"article","venue":"Machine Learning and Knowledge Extraction","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Generative grammar; Search engine indexing; Artificial intelligence; Information retrieval; Machine learning; Natural language processing; Data science","score_opus":0.10801962975602644,"score_gpt":0.4138889138627881,"score_spread":0.3058692841067616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413794792","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08725765,0.002669536,0.7681709,0.0023808398,0.00038406797,0.002319104,0.009175791,0.11412626,0.013515873],"genre_scores_gemma":[0.35193938,0.0007953241,0.6267868,0.001176724,0.00014617268,0.0012146238,0.010257106,0.0017384379,0.0059454027],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99810445,0.0010166836,0.00015142829,0.0002927264,0.0003604613,0.00007414134],"domain_scores_gemma":[0.9940295,0.004434654,0.00026217423,0.0007296685,0.00037322892,0.00017075482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027931507,0.0014416332,0.00049593975,0.00218814,0.0003140015,0.0011881115,0.0015196819,0.0016630627,0.014134541],"category_scores_gemma":[0.013634463,0.00037421813,0.0010211413,0.00087108265,0.00058362575,0.0014356984,0.002461324,0.00084160635,0.00372514],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015656723,0.00062272855,0.0064058383,0.0017048954,0.00032143758,0.003427352,0.0018126997,0.030946888,0.035464305,0.0063010724,0.082361355,0.8290658],"study_design_scores_gemma":[0.0009597381,0.0013849704,0.0060074544,0.0003768716,0.00034843228,0.0056071137,0.0014389462,0.748071,0.06357861,0.04330418,0.12868418,0.00023852833],"about_ca_topic_score_codex":0.0018347332,"about_ca_topic_score_gemma":0.0031538624,"teacher_disagreement_score":0.014134541,"about_ca_system_score_codex":0.00056971645,"about_ca_system_score_gemma":0.0006524527,"threshold_uncertainty_score":0.047284782},"labels":[],"label_agreement":null},{"id":"W4413796994","doi":"10.1515/lingvan-2024-0201","title":"Instance memory models as a general computational framework for exploring language processing: bringing the lexicon to life","year":2025,"lang":"en","type":"article","venue":"Linguistics Vanguard","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Lexicon; Computer science; Cognitive science; Natural language processing; Linguistics; Artificial intelligence; Psychology; Cognitive psychology; Philosophy","score_opus":0.057039101299999395,"score_gpt":0.32504685551013524,"score_spread":0.26800775421013584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413796994","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04727067,0.0011416738,0.9266911,0.0055144853,0.000073135154,0.000053917738,0.00026679167,0.000357148,0.01863107],"genre_scores_gemma":[0.788472,0.0011986422,0.20377944,0.0006211465,0.00026447335,0.00024935574,0.00031537248,0.00018344253,0.0049161273],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934465,0.00036052015,0.000034101326,0.00012355496,0.00007910669,0.0000580435],"domain_scores_gemma":[0.99731857,0.0018157244,0.00014486891,0.0004057989,0.00013432844,0.0001806358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016455114,0.00049325486,0.0008785061,0.0018496694,0.00085348357,0.0049187867,0.002205178,0.0015912183,0.0047215167],"category_scores_gemma":[0.0061271642,0.00051201746,0.0015628992,0.0012802248,0.0041464763,0.011974021,0.0025865606,0.0023459608,0.0006624325],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013970929,0.000014056651,0.0003587876,0.00003653947,0.000021620775,0.00006628758,0.0005392333,0.015078076,0.00033589796,0.97667843,0.00067166134,0.006185378],"study_design_scores_gemma":[0.0000056033996,0.0000075438957,0.000060083177,0.000011495994,0.0000067984115,0.00003769531,0.00007973708,0.07329762,0.000074486044,0.9247548,0.0016565342,0.00000770458],"about_ca_topic_score_codex":0.0024667566,"about_ca_topic_score_gemma":0.0026333078,"teacher_disagreement_score":0.0049187867,"about_ca_system_score_codex":0.001398093,"about_ca_system_score_gemma":0.0008835655,"threshold_uncertainty_score":0.015795052},"labels":[],"label_agreement":null},{"id":"W4413925951","doi":"10.1109/tse.2025.3605442","title":"Towards Explainable Vulnerability Detection With Large Language Models","year":2025,"lang":"en","type":"article","venue":"IEEE Transactions on Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Vulnerability (computing); Data science; Programming language; Software engineering; Natural language processing; Computer security","score_opus":0.00918862300238798,"score_gpt":0.22338367195039308,"score_spread":0.2141950489480051,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413925951","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029745858,0.00080626766,0.9376738,0.0013618723,0.000075198186,0.00017119067,0.0018959249,0.027323835,0.00094608066],"genre_scores_gemma":[0.2944218,0.000492443,0.69129163,0.00090100727,0.00012855537,0.00040123225,0.008626661,0.0014070394,0.0023296925],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99659675,0.0016586152,0.00017305942,0.0009257084,0.00048267806,0.00016307273],"domain_scores_gemma":[0.9851916,0.011303282,0.0008434893,0.001465416,0.00097595464,0.00022029906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034324524,0.0023672124,0.0009023211,0.0034088227,0.00064311834,0.001913844,0.0023934187,0.0022518332,0.0028467688],"category_scores_gemma":[0.022010129,0.0010111107,0.0024115671,0.0015717818,0.00113794,0.004958555,0.0037941278,0.0043412484,0.0021028896],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000505737,0.0003813963,0.016500073,0.0012467272,0.00042583267,0.0010940185,0.0024576995,0.18118477,0.02319397,0.01807875,0.03401882,0.7209122],"study_design_scores_gemma":[0.000037087117,0.0000553896,0.00096457935,0.00006226759,0.00006662579,0.00016100511,0.00021220978,0.9524358,0.005537703,0.033696212,0.006734546,0.000036645775],"about_ca_topic_score_codex":0.0050548757,"about_ca_topic_score_gemma":0.011988171,"teacher_disagreement_score":0.0050548757,"about_ca_system_score_codex":0.0013328749,"about_ca_system_score_gemma":0.0020748284,"threshold_uncertainty_score":0.018152773},"labels":[],"label_agreement":null},{"id":"W4413943190","doi":"10.1002/aaai.70025","title":"Recent advances in finetuning multimodal large language models","year":2025,"lang":"en","type":"article","venue":"AI Magazine","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Institute for Catastrophic Loss Reduction","keywords":"Computer science; Artificial intelligence","score_opus":0.011786876258378312,"score_gpt":0.2828238057033132,"score_spread":0.2710369294449349,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413943190","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039333034,0.015854063,0.9295017,0.0023915896,0.00016460315,0.000048252587,0.000106672094,0.0027471744,0.009852901],"genre_scores_gemma":[0.5881542,0.016649028,0.38596874,0.0010258856,0.0005552937,0.00020784489,0.00046457918,0.0010400682,0.005934345],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99945337,0.0001902159,0.00003991437,0.00016715836,0.00011459967,0.00003478665],"domain_scores_gemma":[0.99781126,0.0013810608,0.00012740844,0.00040989585,0.00020334603,0.0000670066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016620798,0.000664089,0.0006458862,0.00055835565,0.00022975827,0.0012915855,0.0014670576,0.0007042841,0.003014455],"category_scores_gemma":[0.0063810037,0.00042488583,0.0005030509,0.00050188444,0.0008766639,0.002603574,0.001926213,0.0013855894,0.0009409282],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016425052,0.00014305406,0.0012825209,0.0005659236,0.00010492422,0.00008116943,0.00042344752,0.27722657,0.033051994,0.049455658,0.0043499568,0.6331505],"study_design_scores_gemma":[0.000016708434,0.00008673667,0.00051916734,0.0000642846,0.0000504757,0.000071154114,0.00009713166,0.8858096,0.01367833,0.066933036,0.032635096,0.00003835926],"about_ca_topic_score_codex":0.0016042343,"about_ca_topic_score_gemma":0.0014084291,"teacher_disagreement_score":0.003014455,"about_ca_system_score_codex":0.000735632,"about_ca_system_score_gemma":0.00051385997,"threshold_uncertainty_score":0.010084331},"labels":[],"label_agreement":null},{"id":"W4413963629","doi":"10.1016/j.jss.2025.112604","title":"Syntactic multilingual probing of pre-trained language models of code","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Federación Española de Enfermedades Raras; Ministerio de Ciencia, Innovación y Universidades","keywords":"Computer science; Natural language processing; Code (set theory); Linguistics; Artificial intelligence; Programming language; Philosophy","score_opus":0.018899914880622877,"score_gpt":0.27564348552860446,"score_spread":0.2567435706479816,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413963629","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5160882,0.00022314479,0.47379568,0.00039826703,0.00006891863,0.00007046907,0.00045313325,0.004688829,0.0042134356],"genre_scores_gemma":[0.93782216,0.000082859,0.05898313,0.00013601057,0.0000130118715,0.00007058746,0.00097104296,0.00044415865,0.0014771117],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99922717,0.0002770775,0.00003411436,0.00026468257,0.00011522824,0.00008175941],"domain_scores_gemma":[0.9960568,0.0023685351,0.00029572222,0.00072189345,0.0004435478,0.00011349566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012186071,0.0010875241,0.00042491866,0.0005536988,0.00032559008,0.0009957635,0.00090909493,0.00073746545,0.0018825658],"category_scores_gemma":[0.010988295,0.00048119243,0.0008408077,0.00044352276,0.0010578738,0.0033875515,0.0021111309,0.0025363849,0.0007132093],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082414685,0.0002637681,0.02520739,0.00042600697,0.0002665469,0.00059134065,0.0023342876,0.4052631,0.087541826,0.027178153,0.004717793,0.44538563],"study_design_scores_gemma":[0.000010858505,0.00015094626,0.0018716834,0.000019684128,0.00002412451,0.0000896555,0.0001797643,0.96216327,0.02188943,0.012304357,0.0012655992,0.000030618226],"about_ca_topic_score_codex":0.003000492,"about_ca_topic_score_gemma":0.00479357,"teacher_disagreement_score":0.003000492,"about_ca_system_score_codex":0.0008952991,"about_ca_system_score_gemma":0.0010419571,"threshold_uncertainty_score":0.0064959526},"labels":[],"label_agreement":null},{"id":"W4414015822","doi":"10.11159/cist25.141","title":"Improved BoW-BoC Indexing for Short Texts Using Large Language Models","year":2025,"lang":"en","type":"article","venue":"Proceedings of the World Congress on Electrical Engineering and Computer Systems and Science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Search engine indexing; Computer science; Natural language processing; Artificial intelligence; Information retrieval","score_opus":0.01039318354331234,"score_gpt":0.23615430658930603,"score_spread":0.2257611230459937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414015822","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016672239,0.0030284843,0.9692029,0.00034290215,0.00031684997,0.0003038875,0.0012066519,0.006466193,0.0024599305],"genre_scores_gemma":[0.1791125,0.0028973247,0.7910742,0.00055975287,0.00087653275,0.0008238588,0.013040734,0.000861377,0.01075374],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99811256,0.0003536529,0.00016326286,0.00042152664,0.00075697404,0.00019209419],"domain_scores_gemma":[0.9977986,0.00074074813,0.0002171353,0.00044282043,0.00069530593,0.00010538792],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016133676,0.001478442,0.002160471,0.0057622567,0.0009827071,0.0020923375,0.0019634352,0.0013048763,0.0038868275],"category_scores_gemma":[0.006604422,0.00043046064,0.0016340106,0.007882917,0.0006392603,0.005593128,0.0021186178,0.002180655,0.005070416],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036116614,0.00037725124,0.0011040464,0.0005221252,0.00014859521,0.00020557168,0.00032035576,0.023373477,0.044959128,0.019260107,0.029563937,0.8798043],"study_design_scores_gemma":[0.00010120563,0.00032993703,0.0016911642,0.00006476456,0.00014157526,0.00054051,0.0002328461,0.92675817,0.021444913,0.025159866,0.023410197,0.00012477292],"about_ca_topic_score_codex":0.008478318,"about_ca_topic_score_gemma":0.008921156,"teacher_disagreement_score":0.008478318,"about_ca_system_score_codex":0.0009298942,"about_ca_system_score_gemma":0.0020420563,"threshold_uncertainty_score":0.016857922},"labels":[],"label_agreement":null},{"id":"W4414029501","doi":"10.1101/2025.08.31.672925","title":"What Large Language Models Know About Plant Molecular Biology","year":2025,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Nautical Research Society","funders":"Agencia Nacional de Investigación y Desarrollo; Consejo Nacional de Investigaciones Científicas y Técnicas","keywords":"Computational biology; Biology; Computer science","score_opus":0.012640217897755615,"score_gpt":0.24252673282267453,"score_spread":0.2298865149249189,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414029501","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28086597,0.040062353,0.52592117,0.040977176,0.00081462815,0.00038034463,0.06439184,0.012842085,0.033744488],"genre_scores_gemma":[0.7521477,0.01219729,0.14617448,0.0040282873,0.00082204124,0.000619246,0.07896833,0.0014796744,0.003563063],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99329317,0.004050031,0.00027319192,0.0013886684,0.000799651,0.0001953152],"domain_scores_gemma":[0.94535697,0.045244325,0.0014703899,0.0050225817,0.0021346195,0.0007711764],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011154033,0.001717025,0.00123848,0.0038747229,0.00085113244,0.005719197,0.001917364,0.0023193539,0.0042939982],"category_scores_gemma":[0.052672822,0.0008991067,0.0022725512,0.0020546948,0.0010952747,0.014297617,0.0019255384,0.0030004035,0.004132482],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011231666,0.0007419834,0.10308287,0.0068200887,0.0023202146,0.000921347,0.0082352795,0.16197856,0.015527439,0.04432332,0.09208877,0.56283695],"study_design_scores_gemma":[0.0001696854,0.00030004326,0.02639661,0.0023423142,0.0008508555,0.0009653586,0.0028615738,0.57688147,0.0094589135,0.21892484,0.16059186,0.00025649602],"about_ca_topic_score_codex":0.0064044613,"about_ca_topic_score_gemma":0.006517125,"teacher_disagreement_score":0.98884594,"about_ca_system_score_codex":0.0013691253,"about_ca_system_score_gemma":0.0025010982,"threshold_uncertainty_score":0.05898893},"labels":[],"label_agreement":null},{"id":"W4414042672","doi":"10.1007/978-3-032-05607-8_19","title":"Hybrid Ontology Matching for Company Name Alignment: Combining Text Matching, String Similarity, SBERT, and Siamese Networks in Italian and German Job Market Data","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Global Health Research","funders":"","keywords":"Matching (statistics); Discriminative model; Normalization (sociology); Ontology; Robustness (evolution); String searching algorithm; Ontology alignment; Cosine similarity; String (physics); String metric","score_opus":0.023144626615993393,"score_gpt":0.2758699511476817,"score_spread":0.25272532453168833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414042672","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5197698,0.0037454963,0.39447215,0.0013460646,0.00034229035,0.00037840512,0.04673211,0.015669916,0.017543802],"genre_scores_gemma":[0.5972571,0.0011523496,0.28957102,0.00018662048,0.00020930241,0.0002396709,0.103611566,0.0009699408,0.00680242],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983789,0.0005008044,0.00019476205,0.00041840164,0.00034103822,0.00016624601],"domain_scores_gemma":[0.9980332,0.0010845093,0.00014956853,0.00035651337,0.0002979943,0.00007817717],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020085769,0.00053837214,0.00069657643,0.007668624,0.00091417367,0.0018500763,0.0010016174,0.00083312555,0.0029552884],"category_scores_gemma":[0.005265151,0.00021416998,0.0011040458,0.008391941,0.0003559729,0.0033975954,0.0015208777,0.000615718,0.0024819355],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071102194,0.0007060703,0.05785137,0.00060088217,0.0004925603,0.00042033193,0.0006738128,0.026701419,0.01509054,0.0133286165,0.05845299,0.82497036],"study_design_scores_gemma":[0.000115218725,0.00016452855,0.057681803,0.00016213076,0.00040236927,0.00066497736,0.0016559637,0.8273269,0.017272081,0.0419971,0.052443452,0.00011346564],"about_ca_topic_score_codex":0.020728357,"about_ca_topic_score_gemma":0.042473402,"teacher_disagreement_score":0.020728357,"about_ca_system_score_codex":0.0008575946,"about_ca_system_score_gemma":0.0014560409,"threshold_uncertainty_score":0.04121542},"labels":[],"label_agreement":null},{"id":"W4414132264","doi":"10.1017/rsm.2025.10031","title":"StudyTypeTeller—Large language models to automatically classify research study types for systematic reviews","year":2025,"lang":"en","type":"article","venue":"Research Synthesis Methods","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"Universität Zürich; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung","keywords":"Generative grammar; Transformer; Systematic review; Language model; Scientific literature; Encoder","score_opus":0.4744790487850936,"score_gpt":0.5983134327456537,"score_spread":0.12383438396056012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414132264","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028296992,0.019368716,0.6722265,0.012875956,0.0013605958,0.006488914,0.20154929,0.051580723,0.006252355],"genre_scores_gemma":[0.110829,0.0032861682,0.7786722,0.0029557962,0.00032197556,0.0071919244,0.09346777,0.0015825749,0.0016925685],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98314315,0.010529223,0.0026626321,0.0022408285,0.0012131451,0.00021104758],"domain_scores_gemma":[0.83415693,0.15031163,0.0050079655,0.0061121467,0.0037494535,0.0006618756],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.032511175,0.002455671,0.0016228936,0.011218994,0.0011078955,0.0035515353,0.002919576,0.0030083463,0.010621078],"category_scores_gemma":[0.11241241,0.0015440396,0.005689278,0.006012614,0.0010951285,0.0056063053,0.003762526,0.00354038,0.0049548764],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020524561,0.000320289,0.01653356,0.056880694,0.0034930825,0.0015479452,0.0039934753,0.050444692,0.011985123,0.022015572,0.20466739,0.6260657],"study_design_scores_gemma":[0.002143829,0.00068367075,0.008775045,0.01008652,0.004104834,0.0017960968,0.001284694,0.5431052,0.013745965,0.12503213,0.28871536,0.00052663224],"about_ca_topic_score_codex":0.006043514,"about_ca_topic_score_gemma":0.032398004,"teacher_disagreement_score":0.9674888,"about_ca_system_score_codex":0.0033840758,"about_ca_system_score_gemma":0.008725713,"threshold_uncertainty_score":0.17193758},"labels":[],"label_agreement":null},{"id":"W4414153271","doi":"10.1016/j.eswa.2025.129655","title":"The power of text similarity in identifying AI-LLM paraphrased documents: The case of BBC news articles and ChatGPT","year":2025,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Similarity (geometry); Task (project management); Benchmark (surveying); Generative grammar; Power (physics); Revenue","score_opus":0.016926619472284115,"score_gpt":0.29931495088042326,"score_spread":0.2823883314081391,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414153271","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80364007,0.0052229036,0.10819483,0.004864914,0.00027201232,0.0005000415,0.0019819224,0.0017284559,0.07359492],"genre_scores_gemma":[0.9546155,0.00056429807,0.039622147,0.00022137673,0.00020256943,0.00006708888,0.0010875418,0.00026474916,0.0033547946],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9941332,0.003259786,0.0004464957,0.0006202885,0.0011642511,0.00037598063],"domain_scores_gemma":[0.9056955,0.078072175,0.004010832,0.0045287814,0.0068525784,0.0008401277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007110476,0.0005407254,0.0007092745,0.011401982,0.0033447442,0.00597668,0.0013429641,0.003148895,0.0055076075],"category_scores_gemma":[0.068942346,0.0004275022,0.00047387645,0.008124547,0.002157419,0.009066674,0.0030239352,0.0017238756,0.002694533],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005071802,0.00093377545,0.09735958,0.0031447175,0.00046235486,0.016409164,0.046002317,0.016153103,0.06638074,0.052310295,0.02426023,0.671512],"study_design_scores_gemma":[0.0003554186,0.0010639495,0.12545839,0.0009810438,0.00076467474,0.029965913,0.04423346,0.52203393,0.06491574,0.12567522,0.08407681,0.00047546934],"about_ca_topic_score_codex":0.008955238,"about_ca_topic_score_gemma":0.009777203,"teacher_disagreement_score":0.011401982,"about_ca_system_score_codex":0.0012161887,"about_ca_system_score_gemma":0.0011266979,"threshold_uncertainty_score":0.037604272},"labels":[],"label_agreement":null},{"id":"W4414158748","doi":"10.1007/978-981-95-0988-1_1","title":"Evaluating the Behavior of Small Language Models in Answering Binary Questions","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Security token; Binary number; Natural language; Language model; Binary classification; Natural (archaeology)","score_opus":0.11777774152110751,"score_gpt":0.3699668967910038,"score_spread":0.25218915526989627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414158748","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82662344,0.0037637746,0.15103944,0.003476699,0.00033303548,0.00043132654,0.0012702056,0.004418313,0.0086437585],"genre_scores_gemma":[0.9169098,0.000530636,0.075439274,0.00058423,0.00018726445,0.00020887672,0.002557085,0.00046836093,0.0031143946],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9891667,0.00768354,0.0004609504,0.0012733174,0.0010880189,0.0003275718],"domain_scores_gemma":[0.75157446,0.23891902,0.0016890228,0.004101399,0.0021044568,0.0016116482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01844423,0.001649275,0.00177295,0.0015279069,0.0010102575,0.0034253725,0.002520936,0.0038498167,0.0038500216],"category_scores_gemma":[0.09344564,0.0009257402,0.0011392073,0.0014159704,0.0013350721,0.006770864,0.0020118016,0.0035912998,0.0014977476],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0112788165,0.002953261,0.035413984,0.0011873423,0.00088530127,0.00026927493,0.0017438461,0.4903834,0.014342996,0.015046283,0.020790193,0.4057054],"study_design_scores_gemma":[0.000088772475,0.0003003985,0.0009848687,0.000023603036,0.00007058053,0.000036799396,0.00014120528,0.98796856,0.0020808205,0.007852118,0.00043476932,0.000017403172],"about_ca_topic_score_codex":0.009245876,"about_ca_topic_score_gemma":0.009936208,"teacher_disagreement_score":0.01844423,"about_ca_system_score_codex":0.0021808592,"about_ca_system_score_gemma":0.0013808379,"threshold_uncertainty_score":0.09754354},"labels":[],"label_agreement":null},{"id":"W4414160008","doi":"10.1167/tvst.14.9.18","title":"Advancing Question-Answering in Ophthalmology With Retrieval-Augmented Generation: Benchmarking Open-Source and Proprietary Large Language Models","year":2025,"lang":"en","type":"article","venue":"Translational Vision Science & Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":true,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Medical Research Council; Moorfields Eye Hospital NHS Foundation Trust; Retina UK; Moorfields Eye Charity; Sight Research UK; Department of Health and Social Care; Canadian Institute of Steel Construction; UK Research and Innovation; National Institute for Health and Care Research; Amazon Web Services","keywords":"Benchmarking; MEDLINE; Language model; Comprehension","score_opus":0.015463063963740602,"score_gpt":0.31994860134670616,"score_spread":0.3044855373829656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414160008","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56246686,0.016735505,0.29198563,0.004685978,0.0018220413,0.0023041712,0.012883701,0.09059542,0.016520603],"genre_scores_gemma":[0.7110595,0.0015106704,0.23871635,0.002334149,0.00031349424,0.00074621115,0.038360763,0.0013603421,0.0055985693],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99538475,0.0024456098,0.0003009481,0.001037443,0.0006234422,0.00020785004],"domain_scores_gemma":[0.98545086,0.010569504,0.00032462287,0.00189481,0.0013209626,0.00043918914],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008293473,0.0021240108,0.0010464342,0.002096554,0.00067816366,0.0020759306,0.0034544328,0.002606787,0.0038422574],"category_scores_gemma":[0.023105694,0.0005708168,0.0021525822,0.0012393343,0.0009913921,0.003532029,0.0029713858,0.00342021,0.00311093],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016740587,0.0022975402,0.014570276,0.0019469472,0.0008107948,0.00056802767,0.00089820876,0.2117877,0.01249951,0.004194517,0.043568,0.7051844],"study_design_scores_gemma":[0.00048998185,0.0009313585,0.004117712,0.00013495424,0.00023434823,0.00035848044,0.0003536557,0.95888764,0.014341734,0.0073450026,0.012692472,0.00011264951],"about_ca_topic_score_codex":0.016038189,"about_ca_topic_score_gemma":0.021404697,"teacher_disagreement_score":0.016038189,"about_ca_system_score_codex":0.002271077,"about_ca_system_score_gemma":0.0020773534,"threshold_uncertainty_score":0.043860614},"labels":[],"label_agreement":null},{"id":"W4414161315","doi":"10.2196/68707","title":"Performance of Natural Language Processing for Information Extraction From Electronic Health Records Within Cancer: Systematic Review","year":2025,"lang":"en","type":"review","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Aalborg Universitetshospital; Aalborg Universitet","keywords":"Information extraction; Health records; Electronic health record; Natural language; Unstructured data; Biomedical text mining; Text mining; Text processing; Information processing","score_opus":0.019864155168321823,"score_gpt":0.37269172152925883,"score_spread":0.35282756636093704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414161315","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034807557,0.99282324,0.0010888379,0.000515416,0.00013527951,0.00083534507,0.0007529342,0.00003573321,0.0003324553],"genre_scores_gemma":[0.06719124,0.92203397,0.006436594,0.0008975794,0.00018481424,0.0018611914,0.0012269609,0.000032854314,0.000134838],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9627238,0.01495158,0.013290953,0.0027664711,0.00582998,0.00043718328],"domain_scores_gemma":[0.716673,0.24427165,0.020762008,0.0026924799,0.014960718,0.0006401238],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04375592,0.0021773092,0.008867646,0.019811884,0.0009822316,0.0039261584,0.0030939716,0.002086956,0.0024523595],"category_scores_gemma":[0.2172116,0.0013920751,0.017549124,0.014198191,0.0016608681,0.005922467,0.0023107643,0.0015982031,0.00037459307],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003272248,0.000030057194,0.004154813,0.9093565,0.023438606,0.00010282563,0.0003006678,0.0004192294,0.00015070302,0.00018632179,0.0009967444,0.06053639],"study_design_scores_gemma":[0.000292906,0.00049589813,0.010346811,0.78967106,0.18460464,0.0005147345,0.0005089827,0.0009248023,0.000696541,0.00062427344,0.011212959,0.00010650319],"about_ca_topic_score_codex":0.010934816,"about_ca_topic_score_gemma":0.02380796,"teacher_disagreement_score":0.04375592,"about_ca_system_score_codex":0.0055764946,"about_ca_system_score_gemma":0.017501585,"threshold_uncertainty_score":0.23140615},"labels":[],"label_agreement":null},{"id":"W4414255422","doi":"10.1002/cjce.70080","title":"Learning hydrocracking reaction dynamics via neural <scp>ODEs</scp> : A data‐driven, gradient‐interpretable lumped modelling framework","year":2025,"lang":"en","type":"article","venue":"The Canadian Journal of Chemical Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Interpretability; Ode; Extrapolation; Ordinary differential equation; Artificial neural network; Process (computing); Nonlinear system; Complex dynamics","score_opus":0.012091610434469935,"score_gpt":0.21041121615348052,"score_spread":0.19831960571901058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414255422","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10880842,0.00039652063,0.8833569,0.00040215868,0.000077871104,0.000054383534,0.00046578114,0.0007501124,0.005687851],"genre_scores_gemma":[0.9290392,0.00028115377,0.06610089,0.000088896224,0.000026160726,0.00012186159,0.0004351072,0.00010209055,0.003804692],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9999083,0.000021003452,0.000007762079,0.000024136218,0.00002797851,0.000010748087],"domain_scores_gemma":[0.9996427,0.00018800514,0.00004662809,0.00003130319,0.00007551702,0.000015754047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00051282195,0.00054846174,0.00046677195,0.00028684642,0.00016804246,0.0007164099,0.0008117081,0.00072566344,0.0016581623],"category_scores_gemma":[0.0013407584,0.00035444056,0.0006333801,0.00022498498,0.0005220173,0.00074069155,0.00055266195,0.000917532,0.00025412877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000067012425,0.0000073639326,0.00017710846,0.000019804424,0.000007050315,0.000009973896,0.000006896268,0.9941426,0.00092638505,0.0016878322,0.000065219094,0.0029430953],"study_design_scores_gemma":[4.766965e-7,0.0000017981382,0.000016584723,0.0000010218765,6.086774e-7,7.82653e-7,6.425035e-7,0.9994405,0.0001535176,0.00032112037,0.0000619907,8.2977255e-7],"about_ca_topic_score_codex":0.0094946595,"about_ca_topic_score_gemma":0.007247093,"teacher_disagreement_score":0.0094946595,"about_ca_system_score_codex":0.00061643374,"about_ca_system_score_gemma":0.0008343959,"threshold_uncertainty_score":0.018878818},"labels":[],"label_agreement":null},{"id":"W4414344009","doi":"10.4018/979-8-3373-2474-6.ch008","title":"Enhancing Predictive Analysis with Large Language Models in the Digital Innovation World","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"George Brown College","funders":"","keywords":"Generative grammar; Transformative learning; Relevance (law); Software deployment; Deep learning; Artificial neural network; Model-driven architecture","score_opus":0.012728936843463367,"score_gpt":0.23039532396602633,"score_spread":0.21766638712256295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414344009","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010441622,0.00052260765,0.97862077,0.0014779076,0.00007088931,0.0000672983,0.00016034253,0.0025850125,0.0060536764],"genre_scores_gemma":[0.27202985,0.0013275481,0.71641684,0.00070989784,0.00016246235,0.00024878973,0.0007754016,0.0009274991,0.0074016866],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9968539,0.002018346,0.000085076485,0.0003974806,0.00056753826,0.00007769511],"domain_scores_gemma":[0.9738012,0.023516264,0.0003905699,0.0015313657,0.0005877226,0.00017293585],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005916244,0.0009195362,0.00060817937,0.0012122593,0.00064694154,0.004435506,0.0015013139,0.0012583697,0.005018075],"category_scores_gemma":[0.02824966,0.0005174845,0.0010644865,0.0012988762,0.0017335436,0.00794018,0.0033446886,0.003180988,0.0020603794],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018996562,0.00020640584,0.0030731352,0.0006937696,0.00010721845,0.00031656455,0.0037487287,0.17691043,0.010452942,0.28161734,0.010174942,0.5125086],"study_design_scores_gemma":[0.000018403427,0.00007168263,0.00040652885,0.00012144802,0.00003603027,0.00013607982,0.00053426716,0.6629941,0.0062733036,0.30130833,0.028047021,0.000052832376],"about_ca_topic_score_codex":0.0019275821,"about_ca_topic_score_gemma":0.0029580742,"teacher_disagreement_score":0.005916244,"about_ca_system_score_codex":0.0014287034,"about_ca_system_score_gemma":0.0015551479,"threshold_uncertainty_score":0.031288445},"labels":[],"label_agreement":null},{"id":"W4414417426","doi":"10.1007/978-3-032-06004-4_11","title":"AURA: A Multi-modal Medical Agent for Understanding, Reasoning and Annotation","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Interpretability; Toolbox; Modular design; Segmentation; Set (abstract data type); Medical imaging; Relevance (law); Suite; Annotation","score_opus":0.045509895080306666,"score_gpt":0.29349823960367644,"score_spread":0.24798834452336976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414417426","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006191688,0.000851273,0.9383856,0.001394703,0.0003608744,0.00021086427,0.0014991926,0.024896856,0.02620895],"genre_scores_gemma":[0.098363064,0.0009553921,0.8623803,0.0009924653,0.00029905915,0.00039056424,0.0024541952,0.0020841856,0.03208066],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996239,0.00013182456,0.000024988805,0.000074533455,0.00012313722,0.00002165036],"domain_scores_gemma":[0.99937797,0.00036064666,0.00003649196,0.00007780497,0.000069448906,0.00007760319],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009865284,0.0007058097,0.00052079593,0.0006973689,0.0005235865,0.0019951959,0.0013645529,0.0015573788,0.018662157],"category_scores_gemma":[0.0020273996,0.00044874699,0.0006698495,0.0003903124,0.00059609156,0.0019859816,0.0022353309,0.0012391969,0.006161366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011853143,0.00030223263,0.0016466632,0.0011355178,0.0002102059,0.0015911397,0.00119748,0.021102872,0.047205325,0.10235514,0.21836501,0.6037032],"study_design_scores_gemma":[0.00018806485,0.00022886223,0.0012932937,0.00021747344,0.00018302021,0.0030135175,0.0003483209,0.332291,0.032202303,0.14465682,0.48523158,0.00014568743],"about_ca_topic_score_codex":0.00078435795,"about_ca_topic_score_gemma":0.0016217737,"teacher_disagreement_score":0.018662157,"about_ca_system_score_codex":0.00038249552,"about_ca_system_score_gemma":0.00065789936,"threshold_uncertainty_score":0.062431157},"labels":[],"label_agreement":null},{"id":"W4414478919","doi":"10.1016/j.phycom.2025.102857","title":"A survey of domain generalization in AI-enabled semantic communication: Architecture, challenges and future opportunities","year":2025,"lang":"en","type":"article","venue":"Physical Communication","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Generalization; Domain (mathematical analysis); Semantics (computer science); Domain knowledge; Semantic Web","score_opus":0.06782133331123338,"score_gpt":0.2978231013388399,"score_spread":0.2300017680276065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414478919","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0138610825,0.33454633,0.5939806,0.012692005,0.0010701651,0.00027159665,0.00045794758,0.001451319,0.041668992],"genre_scores_gemma":[0.17152047,0.51837456,0.2921486,0.0031504063,0.002079257,0.00044610133,0.00190508,0.00047528153,0.0099002635],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99803,0.0006277916,0.00020446473,0.0003606036,0.00062009634,0.00015695463],"domain_scores_gemma":[0.9956839,0.0024899275,0.00020684929,0.0006372491,0.00077736774,0.00020471215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003646574,0.0010369952,0.0012618101,0.0032159078,0.0010386914,0.0041620857,0.0020994663,0.0014263347,0.003163001],"category_scores_gemma":[0.00704775,0.00067612855,0.0011947898,0.0056344075,0.0018601215,0.010654361,0.0033032245,0.0026094376,0.0015128511],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006037971,0.00012882035,0.0027930415,0.00323304,0.0001075416,0.00023234326,0.0014907289,0.016968807,0.002118447,0.2207805,0.023284342,0.728802],"study_design_scores_gemma":[0.0000144882015,0.00013019016,0.002013553,0.0023880333,0.00010489925,0.001297335,0.0023134106,0.09459443,0.002788239,0.28674936,0.60747415,0.00013196193],"about_ca_topic_score_codex":0.0034972134,"about_ca_topic_score_gemma":0.0027273318,"teacher_disagreement_score":0.0041620857,"about_ca_system_score_codex":0.0017407567,"about_ca_system_score_gemma":0.0023592599,"threshold_uncertainty_score":0.019285142},"labels":[],"label_agreement":null},{"id":"W4414526103","doi":"10.1007/978-3-031-93109-3_28","title":"Integrating NLP with Decision Support Systems for Automated Resume Filtering and Candidate Shortlisting","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Pipeline (software); Process (computing); Selection (genetic algorithm); Automation; Decision support system; Filter (signal processing); Volume (thermodynamics); Support vector machine","score_opus":0.016079364569231064,"score_gpt":0.24842252213396493,"score_spread":0.23234315756473387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414526103","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005357528,0.00044680803,0.95590836,0.0007683543,0.00019688514,0.00023159655,0.0016436059,0.026985833,0.008460998],"genre_scores_gemma":[0.07381024,0.0005120415,0.9132738,0.00036039038,0.00024661233,0.00024003546,0.0036904237,0.0010542461,0.006812156],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99809855,0.00055712054,0.00025544813,0.00044322945,0.0005519742,0.0000936581],"domain_scores_gemma":[0.99102026,0.006548594,0.0003277865,0.0009996594,0.00094422,0.00015948748],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031372984,0.0009341234,0.0012084754,0.002876925,0.00084477843,0.0051317075,0.0020680656,0.0013290839,0.012676648],"category_scores_gemma":[0.011873193,0.0005965525,0.0010992561,0.0028514196,0.0006165477,0.004119197,0.0020863246,0.0018750927,0.00939811],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002553238,0.00030580504,0.001527071,0.0004983342,0.00010306319,0.0003797615,0.0004002132,0.024999924,0.011393438,0.017079988,0.03594256,0.90711457],"study_design_scores_gemma":[0.00008792556,0.00009658658,0.0010156118,0.00015171441,0.00014693591,0.00028182557,0.0003209152,0.7831458,0.023316734,0.1104299,0.080913335,0.00009280372],"about_ca_topic_score_codex":0.0058435197,"about_ca_topic_score_gemma":0.0073926044,"teacher_disagreement_score":0.012676648,"about_ca_system_score_codex":0.001009447,"about_ca_system_score_gemma":0.0014219235,"threshold_uncertainty_score":0.042407632},"labels":[],"label_agreement":null},{"id":"W4414527039","doi":"10.1007/978-3-031-96684-2_12","title":"Can ChatGPT Make Explanatory Inferences? Benchmarks for Abductive Reasoning","year":2025,"lang":"en","type":"book-chapter","venue":"Synthese Library/Synthese library","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Abductive reasoning; Generative grammar; Set (abstract data type); Inference; Verbal reasoning; Generative model; Non-monotonic logic; Visual reasoning","score_opus":0.014534718259256268,"score_gpt":0.214162609239223,"score_spread":0.19962789097996675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414527039","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.063664734,0.0042508994,0.6857688,0.013975761,0.0009906974,0.00041628597,0.0043303575,0.009726008,0.21687652],"genre_scores_gemma":[0.5523938,0.0020689415,0.41041163,0.0009593828,0.00055049855,0.00039482335,0.008618579,0.0026182495,0.021984119],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99381185,0.0030065903,0.0003224876,0.0007593406,0.0017364594,0.0003633053],"domain_scores_gemma":[0.8934144,0.090481065,0.0010659983,0.00981574,0.004213133,0.0010096505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008044717,0.0012137989,0.0009963356,0.0021154457,0.0018141858,0.0068805283,0.004464364,0.0030824787,0.046023577],"category_scores_gemma":[0.09972426,0.0008125744,0.001398075,0.0030845716,0.0030067374,0.019758774,0.004314309,0.004946805,0.008094726],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009031141,0.00033633155,0.0027217926,0.0009235263,0.00010186977,0.0004080978,0.0017112021,0.033695683,0.0013341778,0.5802274,0.05272393,0.3249129],"study_design_scores_gemma":[0.00008108804,0.000047570196,0.0005049511,0.00029746917,0.00005474513,0.00013909377,0.00083276496,0.20129156,0.003590568,0.76864207,0.024487302,0.0000308791],"about_ca_topic_score_codex":0.0034974758,"about_ca_topic_score_gemma":0.0036408945,"teacher_disagreement_score":0.046023577,"about_ca_system_score_codex":0.0016345482,"about_ca_system_score_gemma":0.0019386119,"threshold_uncertainty_score":0.15396422},"labels":[],"label_agreement":null},{"id":"W4414536420","doi":"10.48550/arxiv.2508.13408","title":"NovoMolGen: Rethinking Molecular Language Model Pretraining","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Ministry of Education, India; Indian Institute of Technology Madras; Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Generative grammar; Chemical space; Generative model; Property (philosophy); Scalability; String (physics); Language model; Coherence (philosophical gambling strategy)","score_opus":0.06047665394454849,"score_gpt":0.2974667084531031,"score_spread":0.2369900545085546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414536420","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07234373,0.0019289462,0.8959593,0.0010665486,0.0002804588,0.00020109292,0.002046548,0.018611824,0.0075615738],"genre_scores_gemma":[0.48871368,0.001505412,0.49197587,0.0012022743,0.00010735271,0.00052345404,0.008253373,0.002442239,0.0052762525],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996146,0.00011685238,0.000022039556,0.000116848256,0.000093457835,0.000036234997],"domain_scores_gemma":[0.9981838,0.0012880004,0.000073648764,0.00027235172,0.000120304525,0.000061824016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012444634,0.0009663427,0.0007386615,0.0006463631,0.00035861132,0.0010612645,0.0023488845,0.0011911864,0.005373001],"category_scores_gemma":[0.005473805,0.00058345275,0.001033561,0.0005875567,0.00055081287,0.002455198,0.0013521996,0.0028604676,0.0025989262],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026410085,0.00022907676,0.003096993,0.0007600934,0.0001877105,0.00026226172,0.00019468274,0.66085416,0.014295012,0.03145353,0.018224496,0.27017787],"study_design_scores_gemma":[0.00002941464,0.000058535443,0.00010467015,0.000024342493,0.000016441763,0.000044338096,0.00002126274,0.98157537,0.004921715,0.009398496,0.0037958145,0.00000954066],"about_ca_topic_score_codex":0.0028748084,"about_ca_topic_score_gemma":0.008767461,"teacher_disagreement_score":0.005373001,"about_ca_system_score_codex":0.00093474524,"about_ca_system_score_gemma":0.0015008271,"threshold_uncertainty_score":0.017974496},"labels":[],"label_agreement":null},{"id":"W4414571145","doi":"10.1007/s11633-025-1563-3","title":"A Survey on Personalized Content Synthesis with Diffusion Models","year":2025,"lang":"en","type":"article","venue":"Machine Intelligence Research","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Hong Kong Polytechnic University; Impact Fund","keywords":"Overfitting; Fidelity; Set (abstract data type); Key (lock); Adaptation (eye); Focus (optics)","score_opus":0.23222884324107734,"score_gpt":0.3960968368754752,"score_spread":0.16386799363439786,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414571145","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0139478715,0.08914242,0.8751256,0.0024387015,0.00056792854,0.0002175118,0.0011508657,0.005387886,0.012021208],"genre_scores_gemma":[0.3148293,0.13316406,0.5210487,0.0017919539,0.0024089979,0.0005981468,0.006056728,0.002705922,0.017396154],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99826694,0.00066418375,0.00013208097,0.00050172315,0.00035730388,0.00007779924],"domain_scores_gemma":[0.99552345,0.0033612803,0.00012999223,0.0005659992,0.00034526686,0.00007408084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024735583,0.0014517916,0.001513839,0.002296773,0.00041369078,0.0022216537,0.0019472621,0.001753893,0.0053382535],"category_scores_gemma":[0.011092751,0.0009313857,0.0016132625,0.0025270202,0.0005455311,0.003588982,0.0012673911,0.0014207384,0.0033227075],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014512247,0.00014416076,0.0015354565,0.0021327462,0.00019321404,0.00009808882,0.00031104323,0.06500787,0.003944934,0.017049897,0.019256959,0.8901805],"study_design_scores_gemma":[0.000048700356,0.00023346183,0.002222201,0.0009497848,0.00022425783,0.00052479043,0.00030003092,0.7910468,0.009474282,0.043277003,0.1515934,0.00010534718],"about_ca_topic_score_codex":0.004299314,"about_ca_topic_score_gemma":0.002923607,"teacher_disagreement_score":0.0053382535,"about_ca_system_score_codex":0.00096270314,"about_ca_system_score_gemma":0.00082063675,"threshold_uncertainty_score":0.017858207},"labels":[],"label_agreement":null},{"id":"W4414596705","doi":"10.1162/coli.a.24","title":"Are Formal and Functional Linguistic Mechanisms Dissociated in Language Models?","year":2025,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Azrieli Foundation; Universiteit van Amsterdam; European Commission; Open Philanthropy Project","keywords":"Set (abstract data type); Formal language; Task (project management); Formal system; Theory; Formal methods; Functional approach","score_opus":0.021967691500343927,"score_gpt":0.2635182464301234,"score_spread":0.24155055492977948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414596705","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6523574,0.00057418476,0.3280158,0.002760065,0.00003727314,0.0000704815,0.00022878153,0.0009901022,0.014965886],"genre_scores_gemma":[0.9872378,0.000077751654,0.011771262,0.00011281749,0.000013472778,0.000040096475,0.00011900502,0.000089129455,0.0005388161],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985482,0.0007133334,0.00007650834,0.00032345275,0.00019863897,0.00013976771],"domain_scores_gemma":[0.9933427,0.0032976368,0.00087588717,0.0016784531,0.00040667655,0.0003986343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026526998,0.00044504722,0.0007367321,0.00087459857,0.00045826443,0.0034455915,0.0015527968,0.0010387276,0.003628073],"category_scores_gemma":[0.014960797,0.0006263213,0.00074595184,0.0005182076,0.003406756,0.009005756,0.0028599696,0.0018080496,0.0005069661],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007810328,0.00040134552,0.033328976,0.0006706572,0.00059678534,0.0002973386,0.0053968877,0.1093069,0.05453371,0.6447084,0.0019530532,0.14802492],"study_design_scores_gemma":[0.00006951888,0.0001260648,0.014178028,0.000052354335,0.00008999861,0.0001745281,0.00067222473,0.2524191,0.0058903047,0.7246602,0.0016060346,0.0000616939],"about_ca_topic_score_codex":0.0012547192,"about_ca_topic_score_gemma":0.0010854929,"teacher_disagreement_score":0.003628073,"about_ca_system_score_codex":0.0009596312,"about_ca_system_score_gemma":0.000721395,"threshold_uncertainty_score":0.014029026},"labels":[],"label_agreement":null},{"id":"W4414602263","doi":"10.1057/s41599-025-05834-4","title":"LLMs as annotators: the effect of party cues on labelling decisions by large language models","year":2025,"lang":"en","type":"article","venue":"Humanities and Social Sciences Communications","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Universität Wien","keywords":"Politics; Statement (logic); Test (biology); Motivated reasoning; Replicate","score_opus":0.07625001237173913,"score_gpt":0.3429990491010771,"score_spread":0.266749036729338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414602263","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73537993,0.0025629497,0.22345813,0.0049354727,0.0011522209,0.0015948677,0.0020205036,0.003041473,0.02585449],"genre_scores_gemma":[0.93041813,0.0002587017,0.060258146,0.0028180948,0.00021481226,0.0012289807,0.0014826669,0.000888212,0.0024322635],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.82124525,0.14695911,0.006561264,0.014388354,0.009535736,0.0013103031],"domain_scores_gemma":[0.30736634,0.5910538,0.036641665,0.049576584,0.01279397,0.002567654],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.10631556,0.0020002706,0.0014543945,0.0016041701,0.0024314274,0.0048427754,0.002078812,0.0031109154,0.004189137],"category_scores_gemma":[0.5048245,0.001311302,0.0012299769,0.0019826735,0.0040146303,0.0076333475,0.0069460124,0.0041967803,0.0025991388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01814,0.0011267797,0.36370462,0.004532148,0.004283752,0.0011747418,0.056989525,0.036022633,0.0747926,0.029629834,0.031909946,0.3776935],"study_design_scores_gemma":[0.003174563,0.003621881,0.23725377,0.0020562853,0.0018543329,0.0027110782,0.014124866,0.43264797,0.08759615,0.15679759,0.056253944,0.0019076763],"about_ca_topic_score_codex":0.0053903908,"about_ca_topic_score_gemma":0.008509548,"teacher_disagreement_score":0.89368445,"about_ca_system_score_codex":0.0022055642,"about_ca_system_score_gemma":0.001849814,"threshold_uncertainty_score":0.5622572},"labels":[],"label_agreement":null},{"id":"W4414681735","doi":"10.1007/978-3-032-02725-2_10","title":"Advancing Imminent Fracture Risk Prediction: Integrating Machine Learning with Enhanced Feature Engineering","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Lift (data mining); Feature engineering; Upsampling; Risk assessment; Ensemble learning; Deep learning; Missing data","score_opus":0.003651834518784236,"score_gpt":0.20123297478912347,"score_spread":0.19758114027033924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414681735","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032285098,0.0020233374,0.958091,0.00063284737,0.00018513744,0.00003395216,0.00033520002,0.0018659056,0.0045474456],"genre_scores_gemma":[0.58622485,0.0018315366,0.40511626,0.00023244662,0.00054220157,0.00006118469,0.0010443965,0.0002956257,0.0046515116],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996989,0.000057042707,0.000016234288,0.000078836885,0.00012112089,0.00002803878],"domain_scores_gemma":[0.9987098,0.0007498113,0.000113255424,0.000121962286,0.0002625584,0.00004265165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009030333,0.0007568146,0.001117172,0.0010139897,0.00023518044,0.0013721623,0.001268179,0.0008240945,0.002554997],"category_scores_gemma":[0.004145149,0.00031002946,0.0008126194,0.0008755753,0.00021999532,0.0016727581,0.00101371,0.0011470469,0.0012146848],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000099232246,0.00018018093,0.005625877,0.00012484386,0.00009559649,0.00008216668,0.00003681544,0.1977627,0.0050024916,0.0041923197,0.007800447,0.77899736],"study_design_scores_gemma":[0.000005008205,0.00004155305,0.0010003814,0.000013172822,0.000020725596,0.000046895715,0.000009839893,0.9875794,0.0014092759,0.008471585,0.0013924489,0.000009683645],"about_ca_topic_score_codex":0.0027342387,"about_ca_topic_score_gemma":0.0030831369,"teacher_disagreement_score":0.0027342387,"about_ca_system_score_codex":0.00028303938,"about_ca_system_score_gemma":0.000470461,"threshold_uncertainty_score":0.008547246},"labels":[],"label_agreement":null},{"id":"W4414730095","doi":"10.1007/978-981-96-6929-5_34","title":"Syntax-Constraint-Aware SCABERT: Syntactic Knowledge as a Ground Truth Supervisor of Attention Mechanism via Augmented Lagrange Multipliers","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in networks and systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières; Innovation and Economic Development Trois Rivières; Université du Québec à Montréal","funders":"","keywords":"Interpretability; Merge (version control); ENCODE; Supervisor; Mechanism (biology); Word (group theory); Parsing; Lagrange multiplier","score_opus":0.015762662458433947,"score_gpt":0.2304057376540957,"score_spread":0.21464307519566175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414730095","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00678367,0.00008239671,0.9884289,0.00022928261,0.00004205748,0.00004030651,0.00011656603,0.00084186066,0.0034350615],"genre_scores_gemma":[0.47731167,0.00018348801,0.5092149,0.00022366169,0.000091922324,0.00020235007,0.0005137517,0.0007185069,0.011539772],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99941313,0.0001820836,0.00002516681,0.00018426879,0.00012615354,0.000069186855],"domain_scores_gemma":[0.9989367,0.00051702285,0.00007248065,0.0002485844,0.00015566166,0.00006946434],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012876648,0.00073005725,0.0009785901,0.0005422472,0.0005599684,0.0018809321,0.0028541528,0.0015536568,0.010106627],"category_scores_gemma":[0.004905159,0.0007567021,0.0007286244,0.00059765595,0.0013980623,0.004465026,0.003181071,0.0026129622,0.0014189638],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027998176,0.0001545988,0.0004944443,0.00023836437,0.00009395572,0.00015860648,0.00036975363,0.24498287,0.011178543,0.48900226,0.01346128,0.23958533],"study_design_scores_gemma":[0.000011079032,0.000012404567,0.000064423526,0.000012317147,0.000008767646,0.000014021821,0.00001538937,0.8586829,0.0017988713,0.13779004,0.0015814821,0.000008255417],"about_ca_topic_score_codex":0.0031387797,"about_ca_topic_score_gemma":0.0061338115,"teacher_disagreement_score":0.010106627,"about_ca_system_score_codex":0.0010968642,"about_ca_system_score_gemma":0.0018601913,"threshold_uncertainty_score":0.03381002},"labels":[],"label_agreement":null},{"id":"W4414760587","doi":"10.2196/76661","title":"Beyond Chatbots: Moving Toward Multistep Modular AI Agents in Medical Education","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Education","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Chatbot; Workflow; Modular design; Task (project management); Pipeline (software); Quality (philosophy); Task analysis","score_opus":0.015818902725584765,"score_gpt":0.35511350806207764,"score_spread":0.33929460533649286,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414760587","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04582273,0.0005129649,0.93918204,0.0036921042,0.000095974625,0.00024830663,0.000034184246,0.001920172,0.008491625],"genre_scores_gemma":[0.39563736,0.00032309254,0.59853625,0.0007556986,0.00006252132,0.00034181104,0.00008539964,0.0002061639,0.004051737],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9955368,0.0028773532,0.00019072868,0.00056766695,0.0005209189,0.00030659948],"domain_scores_gemma":[0.9894695,0.00571655,0.00092985865,0.0019587956,0.0007533723,0.001171872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077928444,0.0009452799,0.00037548086,0.0007704201,0.0015131681,0.0040384056,0.0023196395,0.0021028805,0.004777006],"category_scores_gemma":[0.013335954,0.00060830143,0.00078077323,0.00045811306,0.00352299,0.005614989,0.007295954,0.0021672742,0.0015694832],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072124036,0.000981548,0.019320933,0.001431966,0.00023192207,0.0008696427,0.026751434,0.13321085,0.041203193,0.2903622,0.011065037,0.4738501],"study_design_scores_gemma":[0.00016524766,0.0009409402,0.003753171,0.000838357,0.00016795979,0.0005528501,0.0048253178,0.5647172,0.017402273,0.26925972,0.13720326,0.0001738045],"about_ca_topic_score_codex":0.0026852463,"about_ca_topic_score_gemma":0.003286938,"teacher_disagreement_score":0.0077928444,"about_ca_system_score_codex":0.0014849422,"about_ca_system_score_gemma":0.0030141782,"threshold_uncertainty_score":0.041212976},"labels":[],"label_agreement":null},{"id":"W4415002193","doi":"10.1002/smr.70057","title":"Evaluation and Improvement of Test Selection for Large Language Models","year":2025,"lang":"en","type":"article","venue":"Journal of Software Evolution and Process","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Key Research and Development Program of China; Université du Luxembourg","keywords":"Selection (genetic algorithm); Test (biology); Margin (machine learning); Process (computing); Harm; Deep learning; Empirical research; Ground truth","score_opus":0.015548742286457652,"score_gpt":0.3038847352079227,"score_spread":0.28833599292146506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415002193","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7105112,0.0041290754,0.25904208,0.0016108818,0.0005100897,0.0004467353,0.0013005382,0.019605912,0.0028434151],"genre_scores_gemma":[0.93863434,0.00012516954,0.056897555,0.0004016583,0.00007722643,0.00015953615,0.0023962485,0.0003730918,0.00093516073],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99086237,0.004914472,0.00065188215,0.0014799151,0.0016604306,0.00043091894],"domain_scores_gemma":[0.9509532,0.037703604,0.0018823144,0.0030791578,0.0048008286,0.0015810479],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0107199205,0.0019961228,0.0013144104,0.0022414587,0.0005444797,0.0011860044,0.00305389,0.0018747359,0.0018154614],"category_scores_gemma":[0.04633984,0.0005353609,0.0010757482,0.00105975,0.0009744908,0.0019693917,0.0018321894,0.0025753297,0.0008810385],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003701656,0.0016220399,0.05437334,0.0005223356,0.00069919595,0.00053052546,0.00020603632,0.37604234,0.020643864,0.002717192,0.016887136,0.5220543],"study_design_scores_gemma":[0.00008963785,0.00034938168,0.0015818945,0.000013537564,0.000029630466,0.000052360097,0.000028833323,0.99038285,0.006244252,0.0008379133,0.00037620487,0.000013549163],"about_ca_topic_score_codex":0.0036211675,"about_ca_topic_score_gemma":0.0046277493,"teacher_disagreement_score":0.0107199205,"about_ca_system_score_codex":0.0012205649,"about_ca_system_score_gemma":0.0017297107,"threshold_uncertainty_score":0.056693077},"labels":[],"label_agreement":null},{"id":"W4415077466","doi":"10.1007/978-3-032-07502-4_3","title":"Mind the Evaluation Gap: Large Language Models for Structured Data Extraction from Radiology Reports","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Benchmarking; Structured prediction; Process (computing); Information extraction; Topic model; Data extraction","score_opus":0.06755437887015867,"score_gpt":0.3353343658884651,"score_spread":0.2677799870183064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415077466","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008430503,0.0017014637,0.9766815,0.0021647399,0.0002132735,0.00011717973,0.002006171,0.006746438,0.0019387777],"genre_scores_gemma":[0.18896917,0.002401227,0.7840906,0.0010211201,0.0007160557,0.00044965677,0.013140433,0.0032959962,0.0059157982],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960169,0.0022165664,0.00034905158,0.0005274883,0.0007847813,0.00010513895],"domain_scores_gemma":[0.97770894,0.01838822,0.0004974609,0.0018859846,0.001260912,0.00025855575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008505474,0.0010577383,0.0011756787,0.001592117,0.000575359,0.004001783,0.0022015844,0.0013627768,0.0031397466],"category_scores_gemma":[0.026004793,0.00091711077,0.0016470851,0.0019555301,0.0008224835,0.007286891,0.0022361288,0.0034946052,0.0029205852],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060374296,0.0003057907,0.0030650424,0.00094309024,0.00032629992,0.00030638734,0.0012686745,0.04684379,0.009448832,0.054641165,0.08360247,0.7986448],"study_design_scores_gemma":[0.00006560109,0.00011314699,0.0009877797,0.00019650762,0.00015857739,0.00022892251,0.00034300657,0.81515163,0.009472625,0.13898349,0.03423402,0.00006468699],"about_ca_topic_score_codex":0.0025493442,"about_ca_topic_score_gemma":0.0044457684,"teacher_disagreement_score":0.008505474,"about_ca_system_score_codex":0.0009493568,"about_ca_system_score_gemma":0.00205849,"threshold_uncertainty_score":0.044981778},"labels":[],"label_agreement":null},{"id":"W4415077526","doi":"10.1007/978-3-032-07502-4_2","title":"SCOPE: Label Extraction of Stroke Diagnosis from Unstructured Medical Reports Using Retrieval-Augmented Generation","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Foothills Medical Centre; University of Calgary","funders":"","keywords":"Parsing; Scope (computer science); Information extraction; Data extraction; Medical imaging; Code (set theory); Language model; Deep learning","score_opus":0.03351766075588018,"score_gpt":0.29133835355716337,"score_spread":0.25782069280128317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415077526","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049174476,0.0035331182,0.876469,0.0011258919,0.0009802668,0.0012156798,0.0150626935,0.043182578,0.009256165],"genre_scores_gemma":[0.16298814,0.0015673438,0.77509767,0.00052210427,0.00073650444,0.0006426946,0.041273694,0.0019265922,0.015245155],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99913967,0.00016785754,0.00007526136,0.00028550645,0.00024547448,0.000086203516],"domain_scores_gemma":[0.9986873,0.0005331497,0.000091630434,0.00025269398,0.0003705753,0.00006459493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010898686,0.001469571,0.0008748339,0.003800253,0.000644464,0.001701009,0.0011857632,0.001625719,0.007045868],"category_scores_gemma":[0.0024498631,0.00045390447,0.0016168708,0.0018323498,0.00046706464,0.001448363,0.0016495027,0.0008922954,0.006947071],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005217559,0.00024287909,0.002753386,0.00071741437,0.00015425103,0.0006210679,0.00029419723,0.0042870874,0.057105098,0.0018244132,0.06513822,0.8663402],"study_design_scores_gemma":[0.000427157,0.0009446195,0.027551252,0.0004939755,0.0009665475,0.004188569,0.000789074,0.64138764,0.16583501,0.02409816,0.13304621,0.0002718393],"about_ca_topic_score_codex":0.0035735029,"about_ca_topic_score_gemma":0.005721258,"teacher_disagreement_score":0.007045868,"about_ca_system_score_codex":0.0004850527,"about_ca_system_score_gemma":0.0013386409,"threshold_uncertainty_score":0.023570776},"labels":[],"label_agreement":null},{"id":"W4415082293","doi":"10.1142/s1793351x25450023","title":"Extending TriRAG for Advancing Retrieval-Augmented Generation Method with Triple-Based Knowledge Graphs for Improved Question Answering","year":2025,"lang":"en","type":"article","venue":"International Journal of Semantic Computing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Question answering; Knowledge graph; Graph; Semantics (computer science); Language model","score_opus":0.02062060831458576,"score_gpt":0.3493886315228967,"score_spread":0.3287680232083109,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415082293","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016016843,0.00062829757,0.95125556,0.00040892817,0.00012603095,0.00030906915,0.0009171129,0.028304761,0.0020333712],"genre_scores_gemma":[0.24494651,0.00031407684,0.74174935,0.0007801698,0.00011166707,0.000512107,0.005355441,0.0014078357,0.0048228586],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983096,0.0007349151,0.00011616685,0.00039390876,0.00036199688,0.00008338372],"domain_scores_gemma":[0.9972466,0.0014055109,0.00011926738,0.0007451276,0.0004022103,0.00008135773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002426997,0.000993848,0.00091622164,0.002054858,0.00049017335,0.0011673188,0.0015872704,0.0012030933,0.005133753],"category_scores_gemma":[0.008060622,0.00032911633,0.0015958073,0.001174284,0.00068651483,0.0028197607,0.0021335555,0.0014287593,0.0030670755],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038196344,0.00040174983,0.002241704,0.00078944594,0.00019319121,0.00052884367,0.0010564094,0.053245332,0.031808816,0.019427255,0.031786166,0.85813904],"study_design_scores_gemma":[0.000114186114,0.00025859228,0.0006781249,0.00006317441,0.00010198056,0.000447809,0.00027076184,0.9165599,0.020893265,0.032758534,0.027776962,0.00007674652],"about_ca_topic_score_codex":0.0030690928,"about_ca_topic_score_gemma":0.0047813696,"teacher_disagreement_score":0.005133753,"about_ca_system_score_codex":0.000648098,"about_ca_system_score_gemma":0.0011905057,"threshold_uncertainty_score":0.017174125},"labels":[],"label_agreement":null},{"id":"W4415124810","doi":"10.1109/icaide65466.2025.11189651","title":"Structured Memory Mechanisms for Stable Context Representation in Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"","keywords":"Forgetting; Language model; Task (project management); Context (archaeology); Memory model; Context-dependent memory; Episodic memory; Semantics (computer science); Recall","score_opus":0.0235943261662629,"score_gpt":0.2916379206362058,"score_spread":0.2680435944699429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415124810","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025551,0.00021980655,0.9719187,0.0001549219,0.000023130684,0.0000376195,0.00007525747,0.0011586444,0.0008609239],"genre_scores_gemma":[0.7468373,0.0003822944,0.24960138,0.00014985175,0.00005199433,0.000376274,0.00027695624,0.00025925157,0.0020647391],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99957293,0.00014961237,0.000040102263,0.00012954124,0.000065068125,0.000042642292],"domain_scores_gemma":[0.9982157,0.0009596875,0.00012342761,0.0004523923,0.0001800189,0.00006886691],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011757672,0.00076273707,0.0007343709,0.00049014343,0.0005130598,0.0015754671,0.0020226624,0.0008849225,0.0026973344],"category_scores_gemma":[0.0055776797,0.0005930614,0.0009899725,0.00038990704,0.00087043084,0.0048698313,0.0015138194,0.0016973725,0.0006933206],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029727072,0.00013025608,0.001236766,0.00031347122,0.00014388404,0.00021017784,0.0009251945,0.7331589,0.018864088,0.10335097,0.0020953042,0.13927379],"study_design_scores_gemma":[0.00001435734,0.00004146818,0.00009434159,0.000011213576,0.00002312046,0.000024133684,0.000029895084,0.96213436,0.0025449614,0.034203865,0.0008672129,0.000011082901],"about_ca_topic_score_codex":0.0030409489,"about_ca_topic_score_gemma":0.0048522158,"teacher_disagreement_score":0.0030409489,"about_ca_system_score_codex":0.0008418274,"about_ca_system_score_gemma":0.0011658409,"threshold_uncertainty_score":0.009023547},"labels":[],"label_agreement":null},{"id":"W4415155604","doi":"10.1007/978-981-95-3358-9_33","title":"Integrating Information Retrieval and LLMs: A Document Retrieval Chatbot in Education Settings","year":2025,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Chatbot; Relevance (law); Pipeline (software); Search engine indexing; Context (archaeology); Document retrieval; Question answering; Natural language; Automatic indexing","score_opus":0.018849564393143906,"score_gpt":0.2897360571952305,"score_spread":0.2708864928020866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415155604","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14971699,0.004522028,0.6238745,0.005544934,0.0014448185,0.0010306751,0.0017877087,0.05979123,0.15228707],"genre_scores_gemma":[0.40668625,0.0015216172,0.3812926,0.0023408297,0.0009585043,0.0012701667,0.002805431,0.004332764,0.1987918],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987056,0.00080146175,0.000055318524,0.00013839759,0.00020643354,0.00009267724],"domain_scores_gemma":[0.99300283,0.0056012697,0.00012313554,0.0003062668,0.00036205424,0.0006044418],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021687453,0.00083379465,0.000725836,0.0014384818,0.0015014847,0.0029588102,0.001948292,0.0020899454,0.029630804],"category_scores_gemma":[0.0059991297,0.0003716091,0.00035161184,0.0012342811,0.0004717714,0.0046110447,0.0026935749,0.0014414645,0.012469407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088893145,0.0015446651,0.0019008215,0.0012117204,0.00005128886,0.00070475024,0.0069562118,0.0036600246,0.030478459,0.016123474,0.12600636,0.8104733],"study_design_scores_gemma":[0.00056272204,0.0024624686,0.009145561,0.0008435173,0.000234241,0.0025824176,0.009461025,0.21229397,0.047311768,0.04584649,0.6688774,0.00037840204],"about_ca_topic_score_codex":0.001372225,"about_ca_topic_score_gemma":0.0031245025,"teacher_disagreement_score":0.029630804,"about_ca_system_score_codex":0.0008962294,"about_ca_system_score_gemma":0.00094329164,"threshold_uncertainty_score":0.09912491},"labels":[],"label_agreement":null},{"id":"W4415247602","doi":"10.48550/arxiv.2505.03176","title":"seq-JEPA: Autoregressive Predictive Learning of Invariant-Equivariant World Models","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fonds de Recherche du Québec - Santé; Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Institut de Valorisation des Données; Canada Excellence Research Chairs, Government of Canada; Canadian Institute for Advanced Research","keywords":"Embedding; Feature learning; Autoregressive model; Representation (politics); Multi-task learning; Inference; Encoder; Aggregate (composite); Flexibility (engineering); Transformation (genetics)","score_opus":0.06392914997216337,"score_gpt":0.2800418428867269,"score_spread":0.21611269291456353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415247602","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012164001,0.0002157685,0.98253965,0.00017785281,0.000050286977,0.000043734184,0.00016336671,0.0036327953,0.0010124987],"genre_scores_gemma":[0.5994707,0.0004619188,0.38849467,0.0007730429,0.00013663582,0.000285021,0.0025636165,0.00090071425,0.006913706],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993175,0.0001911075,0.000028670594,0.0002696505,0.00012486357,0.000068115165],"domain_scores_gemma":[0.9984395,0.00069115794,0.00014505663,0.0004185669,0.00022999334,0.00007569173],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016069763,0.0013824003,0.0010295317,0.0006089636,0.0003190901,0.0011028931,0.0036855077,0.0011482424,0.0027197564],"category_scores_gemma":[0.004115642,0.0009608374,0.0014470911,0.0007453866,0.00094587496,0.002494107,0.0022846947,0.0038567227,0.0011835294],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000121727615,0.00018375694,0.001295095,0.000099873374,0.00019947502,0.0001308978,0.000107995875,0.7360524,0.004799096,0.014182097,0.0065234657,0.23630413],"study_design_scores_gemma":[0.0000033901533,0.000015294112,0.000055819528,0.0000025247357,0.0000048797724,0.000008528267,0.000002730471,0.9941518,0.00060738705,0.0048827175,0.00026110274,0.000003852686],"about_ca_topic_score_codex":0.004706449,"about_ca_topic_score_gemma":0.007934484,"teacher_disagreement_score":0.004706449,"about_ca_system_score_codex":0.00070408854,"about_ca_system_score_gemma":0.0010553624,"threshold_uncertainty_score":0.009358108},"labels":[],"label_agreement":null},{"id":"W4415259383","doi":"10.1145/3757566","title":"Co-Writing with AI, on Human Terms: Aligning Research with User Demands Across the Writing Process","year":2025,"lang":"en","type":"article","venue":"Proceedings of the ACM on Human-Computer Interaction","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Originality; Agency (philosophy); Process (computing); Generative grammar; Complement (music); Value (mathematics); Control (management)","score_opus":0.09970299942659626,"score_gpt":0.4351658932331937,"score_spread":0.33546289380659744,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415259383","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42084932,0.22861402,0.25570762,0.034094848,0.0012440471,0.003980994,0.00088645663,0.0008593852,0.05376331],"genre_scores_gemma":[0.8030948,0.043860886,0.1431822,0.0026085076,0.00020376753,0.0042892457,0.0003645928,0.00028866404,0.0021073653],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.74723715,0.2056717,0.020419054,0.008078202,0.016881132,0.0017128477],"domain_scores_gemma":[0.35030174,0.58975977,0.019997668,0.018043386,0.019879673,0.0020177744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.17392494,0.00088958105,0.0023458079,0.017551443,0.0047378596,0.018490974,0.00294438,0.002430211,0.002270959],"category_scores_gemma":[0.35339743,0.0012858424,0.0016138854,0.015369489,0.013404517,0.024015767,0.00983962,0.0034886887,0.0005924987],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030336253,0.00015634869,0.01464726,0.04559185,0.0006067716,0.0005638226,0.534316,0.00072565576,0.0034999193,0.031547513,0.0017483367,0.3662932],"study_design_scores_gemma":[0.00030840797,0.0012778875,0.02826679,0.11400651,0.0021662705,0.002646191,0.5639569,0.0029835976,0.014717078,0.09527454,0.17409916,0.00029663448],"about_ca_topic_score_codex":0.0028962328,"about_ca_topic_score_gemma":0.006772352,"teacher_disagreement_score":0.17392494,"about_ca_system_score_codex":0.00666332,"about_ca_system_score_gemma":0.021518705,"threshold_uncertainty_score":0.919814},"labels":[],"label_agreement":null},{"id":"W4415262028","doi":"10.48550/arxiv.2510.12117","title":"Locket: Robust Feature-Locking Technique for Language Models","year":2025,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Government of Ontario","keywords":"Evasion (ethics); Scheme (mathematics); Scalability; Language model; Robustness (evolution); Backdoor; Credential; Chatbot","score_opus":0.08470765334721182,"score_gpt":0.20497739550539826,"score_spread":0.12026974215818644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415262028","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007421113,0.000094441624,0.97670925,0.00035302344,0.000046961348,0.00012992909,0.00020506298,0.01356891,0.001471265],"genre_scores_gemma":[0.62128556,0.00019899533,0.36652917,0.00072169036,0.00011460639,0.0005345714,0.00078439707,0.002286216,0.00754489],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963253,0.0013411988,0.00027589392,0.00057925395,0.0011502149,0.00032816597],"domain_scores_gemma":[0.9883868,0.004886513,0.0008269088,0.0049172263,0.0006952865,0.0002873064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046770563,0.0010130898,0.00077045104,0.0009267199,0.0009415353,0.0019915341,0.0036930232,0.0019630755,0.008684212],"category_scores_gemma":[0.024758868,0.0007538606,0.0016885406,0.00066702865,0.0019940177,0.008632104,0.0056846426,0.0040444722,0.003153032],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011691778,0.00046990337,0.0042035365,0.000542133,0.00023285804,0.00071218493,0.0014898838,0.20530973,0.027163591,0.29071212,0.030194396,0.4378004],"study_design_scores_gemma":[0.000060771046,0.00011592403,0.00014174188,0.00003779796,0.000025845555,0.00021525838,0.00007131846,0.88497967,0.01164639,0.09291209,0.00974786,0.000045329576],"about_ca_topic_score_codex":0.0020005237,"about_ca_topic_score_gemma":0.0025240593,"teacher_disagreement_score":0.008684212,"about_ca_system_score_codex":0.0016150359,"about_ca_system_score_gemma":0.0018720218,"threshold_uncertainty_score":0.029051602},"labels":[],"label_agreement":null},{"id":"W4415297078","doi":"10.1016/j.mlwa.2025.100758","title":"Prompt design for medical question answering with Large Language Models","year":2025,"lang":"en","type":"article","venue":"Machine Learning with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Variety (cybernetics); Workflow; Language model; Tree (set theory); Simple (philosophy); Sonnet","score_opus":0.010040119211515531,"score_gpt":0.27721715884281234,"score_spread":0.2671770396312968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415297078","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032769207,0.0013575344,0.86826926,0.0016500495,0.00023214976,0.00070808805,0.0032386142,0.0896019,0.002173261],"genre_scores_gemma":[0.30613485,0.00046242264,0.6806583,0.0009766818,0.00011404054,0.00083651236,0.0073078056,0.0011409733,0.0023685195],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.997617,0.0012223924,0.00024419883,0.0005326098,0.0002710345,0.000112823895],"domain_scores_gemma":[0.9925287,0.0053476864,0.00028345617,0.0008090024,0.00072605483,0.0003051313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005158338,0.0013694305,0.000948911,0.001004489,0.00045376454,0.0020165078,0.0023388914,0.0017466523,0.00774077],"category_scores_gemma":[0.021382863,0.000539762,0.0016164813,0.0007880782,0.0005548566,0.004284337,0.00231139,0.0027901067,0.003769831],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024688242,0.00086037844,0.009751574,0.0035828238,0.00025100386,0.00090734736,0.0024193483,0.14760265,0.027073296,0.023970254,0.058753263,0.72235924],"study_design_scores_gemma":[0.00030012144,0.00033609633,0.00071772654,0.000102819235,0.00007971963,0.00021294708,0.0004091675,0.9192589,0.011019056,0.042430397,0.025078427,0.000054576845],"about_ca_topic_score_codex":0.0036486208,"about_ca_topic_score_gemma":0.007052116,"teacher_disagreement_score":0.00774077,"about_ca_system_score_codex":0.0013990239,"about_ca_system_score_gemma":0.0023787925,"threshold_uncertainty_score":0.027280211},"labels":[],"label_agreement":null},{"id":"W4415299179","doi":"10.3390/make7040121","title":"Small or Large? Zero-Shot or Finetuned? Guiding Language Model Choice for Specialized Applications in Healthcare","year":2025,"lang":"en","type":"article","venue":"Machine Learning and Knowledge Extraction","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Provincial Health Services Authority; University of British Columbia","funders":"","keywords":"Task (project management); Language model; Selection (genetic algorithm); Exploit; Health care; Language understanding; Task analysis","score_opus":0.07084545954893032,"score_gpt":0.3752476515817091,"score_spread":0.30440219203277874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415299179","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3671596,0.0080051925,0.57939184,0.0065967967,0.00049268827,0.0006141945,0.0015651269,0.024910044,0.011264445],"genre_scores_gemma":[0.83429736,0.0008828363,0.15506376,0.0019375508,0.00009620319,0.0003591032,0.0020165036,0.0010581267,0.004288551],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99820864,0.0009829272,0.000091698996,0.00045533775,0.00012124636,0.00014004261],"domain_scores_gemma":[0.995357,0.0035734852,0.00016587783,0.00037854913,0.00032828285,0.00019681019],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037839455,0.001576215,0.00087113486,0.0005919966,0.000362248,0.0016640486,0.0017611226,0.0014777129,0.0040885494],"category_scores_gemma":[0.016858723,0.00057334604,0.00073872914,0.0003585084,0.00061550125,0.0030040701,0.0017653451,0.0029491659,0.0030302443],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016287207,0.00077864225,0.012296095,0.0011230733,0.00030739076,0.0004102984,0.0009228894,0.16594414,0.038341306,0.0036791912,0.024310961,0.7502574],"study_design_scores_gemma":[0.00019173641,0.00053178094,0.0024885524,0.0002194222,0.00016573016,0.00024861374,0.0005369284,0.9524067,0.022916641,0.011391174,0.0088290945,0.00007356532],"about_ca_topic_score_codex":0.005972823,"about_ca_topic_score_gemma":0.010546065,"teacher_disagreement_score":0.005972823,"about_ca_system_score_codex":0.0008824739,"about_ca_system_score_gemma":0.0019690495,"threshold_uncertainty_score":0.020011604},"labels":[],"label_agreement":null},{"id":"W4415300695","doi":"10.1038/s41746-025-02008-z","title":"When helpfulness backfires: LLMs and the risk of false medical information due to sycophantic behavior","year":2025,"lang":"en","type":"article","venue":"npj Digital Medicine","topic":"Topic Modeling","field":"Computer Science","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Google Research; National Center for Advancing Translational Sciences; European Commission; Harvard Catalyst; National Cancer Institute; National Institutes of Health; Harvard University; Patient-Centered Outcomes Research Institute","keywords":"Helpfulness; Vulnerability (computing); Consistency (knowledge bases); Suspect; Risk assessment; Benchmark (surveying); Baseline (sea); Sophistication","score_opus":0.008338394089495093,"score_gpt":0.2471250392254192,"score_spread":0.2387866451359241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415300695","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8783748,0.001517654,0.10388134,0.005501994,0.00021116254,0.00024218172,0.00087137247,0.006339303,0.003060108],"genre_scores_gemma":[0.97533786,0.0001732988,0.021780265,0.00094217557,0.0000607288,0.00005914526,0.00071110186,0.00017599326,0.00075948914],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99097496,0.0058615133,0.00059947575,0.0013542934,0.0008890812,0.00032063245],"domain_scores_gemma":[0.86756516,0.114313744,0.0051181973,0.008489057,0.003163172,0.0013506758],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018891433,0.0011646304,0.0009437573,0.0009743046,0.00068812683,0.0020608872,0.0013345,0.0023083878,0.0014605339],"category_scores_gemma":[0.100352496,0.00056833454,0.00071629655,0.00050431915,0.0015662079,0.0033052093,0.002129895,0.0037895597,0.0007419727],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0065200254,0.0013454822,0.29908243,0.0016195549,0.0008081608,0.0027359943,0.0092821885,0.24170803,0.030104341,0.0094005205,0.020225894,0.37716743],"study_design_scores_gemma":[0.00032987315,0.001579497,0.032726634,0.0003836721,0.00034536468,0.0017823251,0.0018006779,0.9005942,0.02277718,0.02881924,0.008665741,0.00019567463],"about_ca_topic_score_codex":0.004315872,"about_ca_topic_score_gemma":0.004589023,"teacher_disagreement_score":0.018891433,"about_ca_system_score_codex":0.0010920635,"about_ca_system_score_gemma":0.002000579,"threshold_uncertainty_score":0.09990865},"labels":[],"label_agreement":null},{"id":"W4415301335","doi":"10.18280/isi.300807","title":"Scalable Human Oversight for Aligned Large Language Models: A Hybrid Framework for Intent Fidelity","year":2025,"lang":"","type":"article","venue":"Ingénierie des systèmes d information","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Scalability; Fidelity; Key (lock); Scale (ratio); Scope (computer science); High fidelity","score_opus":0.023828058348013895,"score_gpt":0.27978469483690743,"score_spread":0.25595663648889355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415301335","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015999636,0.00029144142,0.9788785,0.00041569924,0.000032860942,0.00008567079,0.00008649411,0.0035809483,0.0006288282],"genre_scores_gemma":[0.63186145,0.0001753713,0.36427397,0.00049749913,0.000119759934,0.00019600967,0.0004936714,0.0006396135,0.0017426139],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9911658,0.0049680653,0.00036746878,0.0019041145,0.0011649125,0.00042960246],"domain_scores_gemma":[0.9770444,0.012810606,0.002250511,0.0052121673,0.0017967226,0.00088562985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009877136,0.0015477312,0.0014761089,0.0010627958,0.0008268665,0.0022493405,0.0027822452,0.0018320794,0.0020351298],"category_scores_gemma":[0.041376945,0.0007490277,0.0010239327,0.0007124898,0.0029102238,0.0055016796,0.0061574024,0.003896312,0.000811347],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006913577,0.00038354768,0.007932728,0.00030290554,0.00017941733,0.00028399305,0.0012467103,0.57396996,0.015205062,0.0373148,0.00654018,0.35594937],"study_design_scores_gemma":[0.00002722126,0.00007526022,0.00025140485,0.000014021434,0.00001153749,0.000046533525,0.0000454181,0.97327876,0.0019617586,0.02341484,0.0008552098,0.000017998074],"about_ca_topic_score_codex":0.0042682565,"about_ca_topic_score_gemma":0.0049756127,"teacher_disagreement_score":0.009877136,"about_ca_system_score_codex":0.0015846337,"about_ca_system_score_gemma":0.0028795828,"threshold_uncertainty_score":0.0522359},"labels":[],"label_agreement":null},{"id":"W4415312233","doi":"10.2196/63908","title":"Clinical Information Extraction From Notes of Veterans With Lymphoid Malignancies: Natural Language Processing Study","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; National Cancer Institute","keywords":"Veterans Affairs; Information extraction; Pipeline (software); Information system; Natural language; Data extraction; Work (physics)","score_opus":0.016441925406890264,"score_gpt":0.34414397789891465,"score_spread":0.3277020524920244,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415312233","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8885065,0.0022775226,0.0647038,0.002634005,0.00020764179,0.0027622227,0.03368853,0.0025771065,0.0026427864],"genre_scores_gemma":[0.7554902,0.0010232963,0.18047243,0.001197166,0.00017042384,0.0015200422,0.05873793,0.00013861494,0.0012499037],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99687105,0.0009829066,0.0006272995,0.000794411,0.0005744786,0.00014982012],"domain_scores_gemma":[0.97086674,0.022434138,0.002131928,0.0013451037,0.0027827616,0.000439313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00438788,0.0005947503,0.00041809084,0.003139481,0.0006573557,0.001295839,0.0011133602,0.0010391007,0.0012451664],"category_scores_gemma":[0.023560686,0.00022333673,0.00090424117,0.0017695575,0.0005745274,0.0009964197,0.0014557345,0.0011225363,0.00074198475],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015562396,0.0017236602,0.42708692,0.0038117238,0.00039520214,0.009589113,0.0077427654,0.0114669455,0.030540759,0.0017837245,0.033864554,0.47043833],"study_design_scores_gemma":[0.00077997043,0.0015896654,0.49145868,0.0019334411,0.00095558143,0.016063597,0.023017425,0.2724624,0.08805907,0.009882494,0.09331568,0.00048202623],"about_ca_topic_score_codex":0.008452594,"about_ca_topic_score_gemma":0.012153316,"teacher_disagreement_score":0.008452594,"about_ca_system_score_codex":0.001152173,"about_ca_system_score_gemma":0.0027890534,"threshold_uncertainty_score":0.023205638},"labels":[],"label_agreement":null},{"id":"W4415315888","doi":"10.48550/arxiv.2510.06557","title":"The Markovian Thinker: Architecture-Agnostic Linear Scaling of Reasoning","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Samsung; Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research; Microsoft Research","keywords":"Initialization; Reinforcement learning; Context (archaeology); Decoupling (probability); Matching (statistics); Continuation; Constant (computer programming); Markov process","score_opus":0.026058050897031136,"score_gpt":0.2681770662058424,"score_spread":0.24211901530881128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415315888","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24464363,0.000695608,0.725519,0.0022219908,0.0002590051,0.00016836321,0.0005192267,0.012550029,0.013423223],"genre_scores_gemma":[0.86411417,0.0001671613,0.1296125,0.0005012575,0.000050866507,0.00020102023,0.0005091696,0.000713918,0.0041298433],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9990183,0.0003232355,0.000045393048,0.00030084932,0.0001633225,0.00014897312],"domain_scores_gemma":[0.99582183,0.0020588068,0.00028364963,0.0011946323,0.0003741148,0.00026702124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019219111,0.0009856734,0.0006813798,0.00034520187,0.00047780716,0.0014300238,0.0023009486,0.0009504828,0.0074856915],"category_scores_gemma":[0.012003256,0.00077792397,0.0008334821,0.00029078728,0.0012979159,0.0038044835,0.0019854493,0.0041619833,0.0015186867],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012861635,0.000598919,0.0059984038,0.00034355553,0.00017343745,0.00024490088,0.0004225684,0.7361932,0.02086517,0.04964486,0.011301882,0.17292687],"study_design_scores_gemma":[0.00005564901,0.00008829552,0.00032101307,0.000014225977,0.000018823688,0.00002022924,0.000019708703,0.96638274,0.0032821647,0.028725134,0.001059072,0.0000129556665],"about_ca_topic_score_codex":0.0047703027,"about_ca_topic_score_gemma":0.0074770204,"teacher_disagreement_score":0.0074856915,"about_ca_system_score_codex":0.0016507414,"about_ca_system_score_gemma":0.0019675836,"threshold_uncertainty_score":0.025042117},"labels":[],"label_agreement":null},{"id":"W4415330518","doi":"10.1109/iccv51701.2025.02067","title":"Bidirectional Likelihood Estimation with Multi-Modal Large Language Models for Text-Video Retrieval","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Indian Institute of Technology, Patna; Neurosciences Research Foundation","keywords":"Prior probability; Normalization (sociology); Video retrieval; Language model; Source code; Calibration","score_opus":0.015806724776915342,"score_gpt":0.2790875369215376,"score_spread":0.26328081214462223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415330518","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0103085395,0.0015289768,0.9796378,0.00052291475,0.00008620342,0.00012727361,0.00065633946,0.005352418,0.0017795013],"genre_scores_gemma":[0.41609338,0.0017886942,0.55788094,0.0012121716,0.0005228422,0.0007628469,0.006705049,0.0016560863,0.013377947],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99880934,0.00044772506,0.00007616023,0.000306744,0.00025276124,0.00010727253],"domain_scores_gemma":[0.99744546,0.0015970787,0.00017933547,0.00031736796,0.00037332263,0.00008740039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002210954,0.0015340084,0.0012846048,0.001310564,0.0005208664,0.0016008569,0.0022758879,0.001597561,0.0057324586],"category_scores_gemma":[0.012552211,0.00059541495,0.0013799407,0.001586163,0.00065638317,0.0035056998,0.002084965,0.0026893253,0.0052714334],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053457136,0.00037168464,0.0021756934,0.00051282917,0.00022985598,0.0002177842,0.0002861248,0.36190903,0.015952762,0.014885759,0.023928093,0.57899576],"study_design_scores_gemma":[0.000022382512,0.000037561807,0.00021095427,0.000017373679,0.00001868373,0.000048739494,0.000029237517,0.9868841,0.0020879367,0.008580602,0.0020407673,0.00002171095],"about_ca_topic_score_codex":0.012756785,"about_ca_topic_score_gemma":0.01675478,"teacher_disagreement_score":0.012756785,"about_ca_system_score_codex":0.001311388,"about_ca_system_score_gemma":0.0013868364,"threshold_uncertainty_score":0.025365055},"labels":[],"label_agreement":null},{"id":"W4415367566","doi":"10.1109/isit63088.2025.11195330","title":"Leveraging Conditional Mutual Information to Improve Large Language Model Fine-Tuning for Classification","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Mutual information; Process (computing); Adaptability; Language model; Information theory; Deep learning","score_opus":0.02828820071938948,"score_gpt":0.2923358372957938,"score_spread":0.26404763657640434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415367566","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02635112,0.0005487793,0.9681952,0.00041171914,0.000059572925,0.00005702015,0.0001258494,0.0029297122,0.001321132],"genre_scores_gemma":[0.62327594,0.00036375088,0.37058258,0.00070689677,0.00015178621,0.0003009675,0.0010503627,0.0008962568,0.0026714567],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983955,0.0006191789,0.00010897632,0.00038910212,0.000321859,0.00016540065],"domain_scores_gemma":[0.9962805,0.0022821312,0.00027014123,0.00055642554,0.00042195123,0.00018876063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036726051,0.0015961204,0.0013582517,0.001278808,0.0006705722,0.0018062382,0.0023839865,0.0013361075,0.0017125445],"category_scores_gemma":[0.014008948,0.0006108263,0.0014390999,0.000917448,0.0010509689,0.003814465,0.0030829343,0.003588958,0.0011225219],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021828288,0.00028050097,0.0032576313,0.00021968699,0.0002599589,0.00010457572,0.0002561922,0.61101484,0.018492552,0.019604135,0.0055431565,0.34074852],"study_design_scores_gemma":[0.000007873654,0.000027881217,0.000118195916,0.000007521574,0.000011365841,0.000012346779,0.00001384262,0.9870603,0.0031156682,0.009113487,0.0005003353,0.000011097174],"about_ca_topic_score_codex":0.0039498163,"about_ca_topic_score_gemma":0.00806179,"teacher_disagreement_score":0.0039498163,"about_ca_system_score_codex":0.0013780128,"about_ca_system_score_gemma":0.002232432,"threshold_uncertainty_score":0.01942277},"labels":[],"label_agreement":null},{"id":"W4415427950","doi":"10.3233/faia251357","title":"Enhancing Few Shot Named Entity Recognition via Label Semantic Description and Diversity Text","year":2025,"lang":"","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Entity linking; Named-entity recognition; Semantic similarity; Context (archaeology); Representation (politics); Semantics (computer science); Semantic role labeling; Generative grammar","score_opus":0.08487523840846307,"score_gpt":0.27366762772297304,"score_spread":0.18879238931450998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415427950","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009497731,0.00041249907,0.97780836,0.00028071707,0.000063161184,0.000083279956,0.0010496099,0.00935534,0.0014492597],"genre_scores_gemma":[0.18196788,0.00044506238,0.8005824,0.0005520273,0.00011432867,0.000275666,0.009735831,0.0010664255,0.005260486],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99779,0.00074803265,0.00011753406,0.0008044794,0.0004521269,0.00008791334],"domain_scores_gemma":[0.9953585,0.0025798064,0.00025979106,0.0011787033,0.0004883376,0.00013480896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031691394,0.0015150886,0.0012920881,0.002105007,0.0007564574,0.002216569,0.003461074,0.0018259985,0.0036374708],"category_scores_gemma":[0.00895303,0.0007297081,0.0016425726,0.002139924,0.0010086106,0.0069685588,0.003433666,0.0031505523,0.0029406967],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039094477,0.0003784905,0.0041280603,0.0007700443,0.00023037494,0.0005291864,0.000914373,0.10353972,0.02561155,0.03867297,0.033247203,0.7915871],"study_design_scores_gemma":[0.00003802416,0.00008190708,0.00070835976,0.00004311917,0.00006442739,0.00025770025,0.00021484544,0.9078739,0.016955953,0.058449596,0.015247629,0.000064427266],"about_ca_topic_score_codex":0.0025645518,"about_ca_topic_score_gemma":0.0065640067,"teacher_disagreement_score":0.0036374708,"about_ca_system_score_codex":0.0010439154,"about_ca_system_score_gemma":0.0012350519,"threshold_uncertainty_score":0.01676017},"labels":[],"label_agreement":null},{"id":"W4415428016","doi":"10.3233/faia251274","title":"Adversarial Topic-Aware Prompt-Tuning for Cross-Topic Automated Essay Scoring","year":2025,"lang":"","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Adversarial system; Robustness (evolution); Classifier (UML); Focus (optics); Limiting; Training set; Feature (linguistics)","score_opus":0.05388806863981957,"score_gpt":0.3233714527852051,"score_spread":0.26948338414538553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415428016","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09171769,0.0013993195,0.88347757,0.00044373036,0.00037628945,0.00031702212,0.0010557108,0.01733547,0.0038771746],"genre_scores_gemma":[0.7536246,0.00035753727,0.23036319,0.00043038724,0.00030322737,0.0005878176,0.0040522725,0.0007694568,0.009511523],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99695694,0.0013875373,0.00016504298,0.0008635626,0.00044306484,0.00018387275],"domain_scores_gemma":[0.99495065,0.0024499714,0.00040540565,0.00094923266,0.00091323775,0.0003315208],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040477677,0.00151124,0.0012814337,0.0009291375,0.00044241283,0.0011979487,0.0019275085,0.0011291003,0.003028255],"category_scores_gemma":[0.013388214,0.00039414404,0.00061542256,0.0008173438,0.00069600675,0.0020450056,0.0026993568,0.0021930211,0.002505442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089516665,0.00059026625,0.008297039,0.00045349222,0.00018458544,0.00030146298,0.00064652535,0.23836751,0.030393101,0.0064521125,0.028110845,0.6853079],"study_design_scores_gemma":[0.00006667625,0.00014446488,0.0012209166,0.00002049153,0.000019723508,0.00008031347,0.000059500224,0.9794226,0.0071059507,0.007937011,0.0038965123,0.000025773785],"about_ca_topic_score_codex":0.0011200054,"about_ca_topic_score_gemma":0.0018609317,"teacher_disagreement_score":0.0040477677,"about_ca_system_score_codex":0.0007121975,"about_ca_system_score_gemma":0.0010773769,"threshold_uncertainty_score":0.021406889},"labels":[],"label_agreement":null},{"id":"W4415428751","doi":"10.3233/faia250938","title":"An Interpretable Quantum-Inspired Model for Multi-Task Natural Language Understanding","year":2025,"lang":"","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute; Université de Montréal","funders":"","keywords":"Interpretability; Semantics (computer science); Natural language; Artificial neural network; Natural language understanding; Benchmark (surveying); Language model; Deep learning","score_opus":0.10961997875705605,"score_gpt":0.33529427124661815,"score_spread":0.2256742924895621,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415428751","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015454609,0.00019747019,0.98002964,0.0007344552,0.00005083133,0.00003103099,0.00012277899,0.00013113717,0.003247957],"genre_scores_gemma":[0.7388214,0.00052367017,0.2525219,0.00036182953,0.00014290564,0.0003039738,0.00031261903,0.000115485745,0.006896216],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996393,0.00015720232,0.0000149430025,0.000093409704,0.000064450156,0.000030698735],"domain_scores_gemma":[0.99910873,0.0005607573,0.00010134137,0.000097809985,0.00008718186,0.00004410542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010245787,0.0004649404,0.0004998275,0.00040569194,0.0003703596,0.0011902086,0.0014296004,0.0011243691,0.0024514033],"category_scores_gemma":[0.003040732,0.00029045064,0.0006532998,0.0005474683,0.00094889384,0.002640721,0.0010351214,0.0018106299,0.00033741572],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000087135806,0.00009333449,0.001015685,0.00017920298,0.00009158644,0.00014131406,0.00045242728,0.49487492,0.0068346905,0.43561298,0.0032077648,0.05740902],"study_design_scores_gemma":[0.0000035884589,0.00001102956,0.00010242979,0.000004797462,0.000005717217,0.00001345814,0.000008819585,0.9077522,0.00023497762,0.09129403,0.00056417333,0.000004874852],"about_ca_topic_score_codex":0.001602777,"about_ca_topic_score_gemma":0.0020328765,"teacher_disagreement_score":0.0024514033,"about_ca_system_score_codex":0.00090072904,"about_ca_system_score_gemma":0.00072608574,"threshold_uncertainty_score":0.008200705},"labels":[],"label_agreement":null},{"id":"W4415442763","doi":"10.2139/ssrn.5637851","title":"Small Language Graph: Large-Scale Model Intelligence on a Small GPU","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Adaptation (eye); Inference; Language model; Domain (mathematical analysis); Isolation (microbiology); Graph; Domain adaptation; Node (physics)","score_opus":0.028497601089149092,"score_gpt":0.2678174905014295,"score_spread":0.2393198894122804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415442763","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04293295,0.00034660354,0.931541,0.0012857047,0.00022649541,0.00009342695,0.0011363879,0.012292975,0.0101444665],"genre_scores_gemma":[0.44861105,0.00041840313,0.5363007,0.00047004566,0.00016313132,0.00021434264,0.002238339,0.0025722613,0.009011666],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99955815,0.00012090339,0.000019565046,0.0001422806,0.00010973841,0.0000493594],"domain_scores_gemma":[0.9983485,0.0007326365,0.000059215483,0.00050744734,0.00020167843,0.00015043402],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005376506,0.00080522115,0.0010688927,0.0008703446,0.00084871985,0.0020561307,0.0021569568,0.0011740648,0.011717811],"category_scores_gemma":[0.004384118,0.0006063851,0.0011527973,0.0015735338,0.00072256743,0.003762871,0.0016625071,0.0022338394,0.0026999717],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005476922,0.00027158455,0.0029889164,0.0004255109,0.00023601255,0.0004445323,0.00057088875,0.5142642,0.01564897,0.1273905,0.07590665,0.26130453],"study_design_scores_gemma":[0.000021723725,0.000012308262,0.00009536351,0.0000038696917,0.000012014144,0.000017970193,0.000036168818,0.94456774,0.0009399277,0.050767414,0.003519447,0.0000060052816],"about_ca_topic_score_codex":0.011776655,"about_ca_topic_score_gemma":0.022610852,"teacher_disagreement_score":0.011776655,"about_ca_system_score_codex":0.0009951949,"about_ca_system_score_gemma":0.0014741556,"threshold_uncertainty_score":0.039200008},"labels":[],"label_agreement":null},{"id":"W4415451240","doi":"10.3390/electronics14214134","title":"Integrating Unstructured EHR Data Using an FHIR-Based System: A Case Study with Problem List Data and an FHIR IPS Model","year":2025,"lang":"en","type":"article","venue":"Electronics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Interoperability; Pipeline (software); Key (lock); Semantic interoperability; Resource (disambiguation); Unstructured data; Component (thermodynamics); Semantic mapping","score_opus":0.06418814022123576,"score_gpt":0.32654962225917816,"score_spread":0.2623614820379424,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415451240","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8277574,0.000607849,0.14734392,0.0050634206,0.00011003337,0.0016277302,0.0047789826,0.002635525,0.010075131],"genre_scores_gemma":[0.76411504,0.00048904214,0.22445296,0.00070483546,0.0000542127,0.00047873918,0.005265637,0.0003398397,0.0040996503],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935375,0.0037442786,0.0006028089,0.00065061706,0.0011963982,0.00026845417],"domain_scores_gemma":[0.98311883,0.012562598,0.00063964847,0.0014831843,0.0017612737,0.00043450674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067957994,0.0005639655,0.0005695536,0.00236248,0.0018524308,0.0031177297,0.0015868049,0.0026238458,0.001608523],"category_scores_gemma":[0.017625794,0.0003460104,0.0010107529,0.0035501383,0.0012189529,0.0032844439,0.0018988142,0.0011745442,0.0007761192],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002412672,0.004789385,0.16172406,0.0037265313,0.0004591199,0.12619309,0.08467777,0.07973641,0.034936015,0.03669056,0.032648504,0.43200582],"study_design_scores_gemma":[0.0005284351,0.0018252315,0.073254555,0.0010090254,0.00059211795,0.036343366,0.07122572,0.49355337,0.07369319,0.021859812,0.22555932,0.00055581436],"about_ca_topic_score_codex":0.022030924,"about_ca_topic_score_gemma":0.024493897,"teacher_disagreement_score":0.022030924,"about_ca_system_score_codex":0.0029897501,"about_ca_system_score_gemma":0.0021735097,"threshold_uncertainty_score":0.04380536},"labels":[],"label_agreement":null},{"id":"W4415473192","doi":"10.2196/78332","title":"Enabling Just-in-Time Clinical Oncology Analysis With Large Language Models: Feasibility and Validation Study Using Unstructured Synthetic Data","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Proof of concept; Clinical Oncology; Unstructured data; Synthetic data; Software; Patient data; Data format","score_opus":0.10765232474025647,"score_gpt":0.42194807222790515,"score_spread":0.31429574748764866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415473192","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.83511645,0.0008164499,0.13453546,0.00192291,0.00022110122,0.0015356932,0.013322207,0.00939988,0.0031298292],"genre_scores_gemma":[0.77405614,0.0002367247,0.19963306,0.0004716729,0.00007289976,0.0010497762,0.023250166,0.00021957002,0.0010100158],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957813,0.002676774,0.0003157137,0.00067062903,0.00041893616,0.00013659477],"domain_scores_gemma":[0.97000587,0.024522722,0.000913433,0.0024994335,0.0015432303,0.000515349],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007167677,0.00096869,0.00040074115,0.0010382244,0.00041337305,0.0013208677,0.0015053378,0.00093491137,0.0016676515],"category_scores_gemma":[0.028514033,0.00027749178,0.0009533705,0.000796347,0.00069406896,0.0014228644,0.0017759614,0.0009844069,0.0009031925],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0048586386,0.0040912805,0.118695505,0.003260079,0.0008467912,0.0028949047,0.005520569,0.31728938,0.029448692,0.008384476,0.04737625,0.45733342],"study_design_scores_gemma":[0.0004260232,0.0013005736,0.01819462,0.00016668772,0.00011931373,0.000785277,0.001656173,0.9322947,0.021554053,0.007044829,0.016362388,0.00009536608],"about_ca_topic_score_codex":0.003569427,"about_ca_topic_score_gemma":0.0041221376,"teacher_disagreement_score":0.007167677,"about_ca_system_score_codex":0.0008939089,"about_ca_system_score_gemma":0.0015642256,"threshold_uncertainty_score":0.037906766},"labels":[],"label_agreement":null},{"id":"W4415474980","doi":"10.1681/asn.2025feyefqjk","title":"Assessing Open-Weight, Large Language Model for Symptom Extraction from Dialysis Notes: A Cost-Effective and Privacy-Preserving Approach","year":2025,"lang":"en","type":"article","venue":"Journal of the American Society of Nephrology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre Hospitalier de l’Université de Montréal; McGill University Health Centre","funders":"","keywords":"Dialysis; Kidney disease; Hemodialysis; MEDLINE; Data extraction","score_opus":0.03270521271704304,"score_gpt":0.3448732067205751,"score_spread":0.31216799400353207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415474980","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36415008,0.0014663318,0.58085996,0.0016323968,0.00028341467,0.0011142595,0.00810805,0.03894853,0.003436891],"genre_scores_gemma":[0.599962,0.00033796957,0.38642457,0.00045455623,0.00008871742,0.000580248,0.010010018,0.0006031421,0.0015388251],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964245,0.001639948,0.00037165027,0.000764761,0.00068427064,0.00011483066],"domain_scores_gemma":[0.98093516,0.01430695,0.0009471439,0.0016529474,0.0018529256,0.0003048363],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053815194,0.0015021554,0.00070582045,0.0013340526,0.00039779817,0.0022636824,0.0013869136,0.0011341594,0.00201186],"category_scores_gemma":[0.025216812,0.0004172261,0.0012701687,0.0006193497,0.0003623221,0.0027055999,0.0017804117,0.0013327821,0.0015926809],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004388544,0.0019149232,0.04754728,0.0022019106,0.0008871973,0.0010039659,0.0014414291,0.10860806,0.03977884,0.0030317083,0.021390991,0.7678051],"study_design_scores_gemma":[0.00027353008,0.0009110099,0.009849652,0.00019203096,0.00037811676,0.000550154,0.0006773783,0.92747366,0.039771173,0.006917817,0.0128831435,0.0001224397],"about_ca_topic_score_codex":0.0033481538,"about_ca_topic_score_gemma":0.004285741,"teacher_disagreement_score":0.0053815194,"about_ca_system_score_codex":0.0010974256,"about_ca_system_score_gemma":0.0017268823,"threshold_uncertainty_score":0.028460562},"labels":[],"label_agreement":null},{"id":"W4415478691","doi":"10.1016/j.jisa.2025.104284","title":"Security concerns for Large Language Models: A survey","year":2025,"lang":"en","type":"article","venue":"Journal of Information Security and Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Competitor analysis; Key (lock); Focus (optics); Security through obscurity; Natural language; Open research; Critical security studies","score_opus":0.02257318162330452,"score_gpt":0.3026937380893925,"score_spread":0.28012055646608797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415478691","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032513145,0.2586567,0.6615614,0.026157673,0.0006537264,0.00026849943,0.0012821882,0.0027008925,0.01620575],"genre_scores_gemma":[0.45296186,0.3094477,0.21081169,0.006614564,0.0066857026,0.0005770987,0.004449575,0.0021314425,0.006320362],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9851362,0.0061214566,0.0015759185,0.001765466,0.0048840097,0.0005169717],"domain_scores_gemma":[0.8507078,0.13223949,0.0036296928,0.007314567,0.0053269803,0.0007814793],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01769172,0.0012780062,0.0027699908,0.004043491,0.0011884096,0.007105273,0.003460974,0.0036845587,0.00421506],"category_scores_gemma":[0.07572624,0.002221791,0.0024571456,0.0051982943,0.0025345145,0.021108149,0.003654127,0.005362556,0.0018957151],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003853786,0.00044887202,0.013197821,0.005371206,0.0005142771,0.0005803381,0.0016441667,0.04072285,0.0026016044,0.321441,0.029735947,0.58335656],"study_design_scores_gemma":[0.00008931855,0.00025360254,0.0030915882,0.0018502048,0.00043805953,0.0026313544,0.0012212552,0.23156224,0.005307013,0.56976116,0.18362226,0.0001719199],"about_ca_topic_score_codex":0.0020081229,"about_ca_topic_score_gemma":0.0014228522,"teacher_disagreement_score":0.01769172,"about_ca_system_score_codex":0.0026916207,"about_ca_system_score_gemma":0.0030250652,"threshold_uncertainty_score":0.093563855},"labels":[],"label_agreement":null},{"id":"W4415499254","doi":"10.54195/irrj.22626","title":"Effectiveness of In-Context Learning for Due Diligence","year":2025,"lang":"en","type":"article","venue":"Information Retrieval Research","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Universiteit van Amsterdam; Canadian Institute of Steel Construction","keywords":"Due diligence; Dependency (UML); Python (programming language); Diligence; Task (project management); Natural language","score_opus":0.048900071996636006,"score_gpt":0.38375496990580404,"score_spread":0.33485489790916806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415499254","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2412555,0.008216811,0.7157658,0.0037832523,0.000658663,0.0005434104,0.0007716794,0.010749779,0.01825514],"genre_scores_gemma":[0.9022862,0.0006109194,0.08897599,0.00060347904,0.00025201845,0.0001482627,0.001235589,0.00060779933,0.005279867],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99067104,0.0045136176,0.000466565,0.0023358949,0.0014819722,0.0005308167],"domain_scores_gemma":[0.95846146,0.027885666,0.0017683787,0.007398293,0.0030831073,0.0014030525],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012983353,0.0016651017,0.0019960068,0.0020794591,0.00250744,0.0029874712,0.0031725306,0.0022006994,0.004642055],"category_scores_gemma":[0.072939634,0.00068550865,0.0013608685,0.0012947965,0.0018321029,0.008217785,0.0034464856,0.0045268824,0.0018618284],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012118535,0.001005711,0.026723126,0.0005696972,0.0003263836,0.0004436578,0.0018263598,0.1910476,0.0069940733,0.02621812,0.021372449,0.72226095],"study_design_scores_gemma":[0.0000897169,0.00045796065,0.0051008025,0.00009507784,0.000106293825,0.00030711532,0.00039276,0.92866886,0.0075032027,0.048657484,0.008527126,0.00009357809],"about_ca_topic_score_codex":0.009666372,"about_ca_topic_score_gemma":0.010628611,"teacher_disagreement_score":0.012983353,"about_ca_system_score_codex":0.0030298394,"about_ca_system_score_gemma":0.0038494838,"threshold_uncertainty_score":0.06866336},"labels":[],"label_agreement":null},{"id":"W4415595662","doi":"10.1016/j.ipm.2025.104449","title":"Leveraging historical information to boost retrieval-augmented generation in conversations","year":2025,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Amazon","keywords":"Boosting (machine learning); Context (archaeology); Text generation; Information system; Information needs","score_opus":0.020043626248497374,"score_gpt":0.2488657579690393,"score_spread":0.22882213172054192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415595662","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21500596,0.0074383584,0.7400053,0.0020587628,0.0006719987,0.0006125706,0.001591759,0.023448288,0.0091669755],"genre_scores_gemma":[0.7883257,0.00076585944,0.19988362,0.0008047422,0.00047331708,0.00032853198,0.0033179,0.0007850639,0.005315188],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980136,0.0008771616,0.000086182765,0.00055169046,0.0002632539,0.00020809592],"domain_scores_gemma":[0.9928681,0.0051107514,0.00021019756,0.0008547804,0.00066274795,0.00029335983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037444755,0.0021121316,0.0016010276,0.0017330627,0.0010240922,0.0016885294,0.002166729,0.0018907405,0.0036938514],"category_scores_gemma":[0.015316476,0.0007297658,0.0010200089,0.0010920458,0.0006435859,0.0036048838,0.0022928307,0.002487734,0.0027820214],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001664143,0.00089709414,0.013365607,0.0007172593,0.00031380943,0.000411883,0.001533674,0.111805044,0.030157274,0.0034580224,0.019545756,0.8161304],"study_design_scores_gemma":[0.00011304356,0.00032648296,0.0026777156,0.000058678896,0.00016336436,0.00027358797,0.00033801878,0.97042894,0.012583674,0.006911816,0.00605003,0.00007465597],"about_ca_topic_score_codex":0.005181458,"about_ca_topic_score_gemma":0.010658086,"teacher_disagreement_score":0.005181458,"about_ca_system_score_codex":0.0007391302,"about_ca_system_score_gemma":0.0014371058,"threshold_uncertainty_score":0.019802868},"labels":[],"label_agreement":null},{"id":"W4415622293","doi":"10.1145/3773285","title":"TaskEval: Assessing Difficulty of Code Generation Tasks for Large Language Models","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Huawei Technologies (Canada)","funders":"","keywords":"Benchmarking; Benchmark (surveying); Task (project management); Code (set theory); Program comprehension; Code generation; Source code","score_opus":0.11296114902424609,"score_gpt":0.36077962888675685,"score_spread":0.24781847986251077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415622293","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7226199,0.0011954497,0.23179214,0.00048511967,0.0002981583,0.002275503,0.010164919,0.021549778,0.009619069],"genre_scores_gemma":[0.78418124,0.0002577919,0.18926722,0.00022270023,0.00007169814,0.0030236393,0.018550184,0.0018983628,0.0025271126],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9846506,0.008590813,0.0018254286,0.0017252542,0.00280406,0.00040372048],"domain_scores_gemma":[0.8590377,0.11224894,0.00992821,0.00926369,0.007584082,0.0019374862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013125193,0.0020994502,0.000896756,0.0041150795,0.00059396564,0.002476002,0.0017139455,0.0016497934,0.003131385],"category_scores_gemma":[0.114503175,0.00047938357,0.0012398509,0.002312798,0.0007499501,0.0038779061,0.0036891461,0.0016134037,0.0016225877],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004863173,0.0037824006,0.16912688,0.006162766,0.0011281398,0.0004591838,0.01101167,0.05393648,0.02888154,0.0077656386,0.058356963,0.65452516],"study_design_scores_gemma":[0.001221494,0.009038024,0.30812398,0.0008278027,0.00051157543,0.00140002,0.0065490142,0.5203095,0.06526015,0.029308772,0.056605384,0.0008443152],"about_ca_topic_score_codex":0.0024732936,"about_ca_topic_score_gemma":0.0033085255,"teacher_disagreement_score":0.013125193,"about_ca_system_score_codex":0.0010064418,"about_ca_system_score_gemma":0.0013353772,"threshold_uncertainty_score":0.06941348},"labels":[],"label_agreement":null},{"id":"W4415683709","doi":"10.3390/info16110937","title":"DefAn: Definitive Answer Dataset for LLM Hallucination Evaluation","year":2025,"lang":"en","type":"article","venue":"Information","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"King Fahd University of Petroleum and Minerals","keywords":"Benchmark (surveying); Benchmarking; Consistency (knowledge bases); Scope (computer science); Scale (ratio); Generative grammar","score_opus":0.03434413762232726,"score_gpt":0.3188867060780858,"score_spread":0.28454256845575854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415683709","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10982383,0.007917889,0.043905634,0.0037240263,0.0018423949,0.0035688467,0.7105834,0.08457306,0.034060907],"genre_scores_gemma":[0.072563246,0.000614227,0.046717294,0.0012175088,0.0001417747,0.0019539448,0.8673287,0.0008680713,0.008595258],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99548006,0.001837879,0.0005125175,0.0008565996,0.0011000737,0.00021287819],"domain_scores_gemma":[0.99199486,0.0039715148,0.00041885613,0.0016833195,0.001517246,0.0004142243],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036120536,0.003385186,0.0010132804,0.0018635687,0.0010183685,0.0014586861,0.0033524083,0.003268928,0.01785914],"category_scores_gemma":[0.020047404,0.00042662694,0.0013987215,0.0013106888,0.0010095581,0.0031580732,0.003277986,0.0032310283,0.015120087],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014590687,0.0009593742,0.0050774347,0.0032648847,0.00017786302,0.0005680969,0.00063985056,0.00797874,0.00487183,0.0028920285,0.85012114,0.121989645],"study_design_scores_gemma":[0.001960077,0.0017948697,0.017392004,0.0009000345,0.00019076046,0.0016373917,0.002321475,0.14153567,0.025339298,0.013350659,0.7932532,0.00032468734],"about_ca_topic_score_codex":0.0079481695,"about_ca_topic_score_gemma":0.015922336,"teacher_disagreement_score":0.01785914,"about_ca_system_score_codex":0.0016143756,"about_ca_system_score_gemma":0.0019939202,"threshold_uncertainty_score":0.059744716},"labels":[],"label_agreement":null},{"id":"W4415706004","doi":"10.36227/techrxiv.176184397.77169395/v1","title":"Security Knowledge Dilution in Large Language Models: How Irrelevant Context Degrades Critical Domain Expertise","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute","funders":"","keywords":"Context (archaeology); Phenomenon; Domain (mathematical analysis); Domain knowledge; Knowledge acquisition; Term (time); Relevance (law)","score_opus":0.02692883516655278,"score_gpt":0.3072965362562432,"score_spread":0.28036770108969045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415706004","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9705396,0.00018506798,0.026745984,0.00033486626,0.000024462208,0.00005403112,0.000084850544,0.0007644721,0.0012667766],"genre_scores_gemma":[0.9917521,0.00005014765,0.007554102,0.000114340495,0.00001287472,0.00003271093,0.000119627424,0.000095681775,0.00026852946],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964593,0.0019434846,0.00023479505,0.00063529506,0.00051259063,0.0002144985],"domain_scores_gemma":[0.9627305,0.030620025,0.0014486909,0.0035405543,0.000809064,0.000851061],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037705598,0.00071610714,0.00059787877,0.00040472485,0.00059306336,0.0015454112,0.0011621442,0.0010880677,0.0015602064],"category_scores_gemma":[0.059521895,0.00059720414,0.000372118,0.00025801332,0.001289598,0.0035312304,0.002609944,0.0018127477,0.0004241475],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005566868,0.0019540444,0.07736301,0.0009916914,0.00051331654,0.0017500973,0.012349476,0.33159286,0.3116929,0.006456265,0.004150523,0.24561907],"study_design_scores_gemma":[0.0002333376,0.0023694546,0.022374881,0.00009509411,0.0002691884,0.00077070214,0.0018064296,0.7972517,0.15185288,0.017720707,0.00508604,0.00016964294],"about_ca_topic_score_codex":0.0019247184,"about_ca_topic_score_gemma":0.0020936541,"teacher_disagreement_score":0.0037705598,"about_ca_system_score_codex":0.00065545517,"about_ca_system_score_gemma":0.00067251525,"threshold_uncertainty_score":0.019940853},"labels":[],"label_agreement":null},{"id":"W4415749046","doi":"10.2196/75932","title":"A Multiagent Summarization and Auto-Evaluation Framework for Medical Text: Development and Evaluation Study","year":2025,"lang":"en","type":"article","venue":"JMIR AI","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; Cleveland Clinic","keywords":"Automatic summarization; Adaptability; Key (lock); Dependency (UML); Scalability; Salient","score_opus":0.04394454882875393,"score_gpt":0.3929848069328207,"score_spread":0.34904025810406675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415749046","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22077188,0.003995159,0.7247828,0.0010925147,0.0003285797,0.005189935,0.0016043342,0.03416481,0.008069908],"genre_scores_gemma":[0.36929476,0.0006552628,0.6213665,0.0002773967,0.00009121255,0.0015300399,0.0030557322,0.0004395978,0.0032894928],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953988,0.0025636435,0.000334377,0.0004788035,0.0010855275,0.00013884684],"domain_scores_gemma":[0.9894757,0.005917603,0.0005641454,0.00063704257,0.0029468855,0.00045860917],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0098526515,0.0013177802,0.00083049043,0.0014999991,0.0005052766,0.0014523413,0.0020116013,0.0012378196,0.0038287952],"category_scores_gemma":[0.016052479,0.00037240563,0.00081901456,0.0007320606,0.00044781997,0.0019537618,0.001183569,0.0012961321,0.0011765368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029879722,0.003011505,0.008228617,0.0024809223,0.00085282477,0.0005432708,0.0015118858,0.121773235,0.028529124,0.00462729,0.02035703,0.8050963],"study_design_scores_gemma":[0.00038120395,0.0020016057,0.0041467985,0.00013213982,0.00026133435,0.00020121627,0.0002729616,0.9674982,0.0137625765,0.0014505046,0.009821441,0.00007002616],"about_ca_topic_score_codex":0.009059133,"about_ca_topic_score_gemma":0.008049098,"teacher_disagreement_score":0.99014735,"about_ca_system_score_codex":0.0019240975,"about_ca_system_score_gemma":0.0018427969,"threshold_uncertainty_score":0.05210638},"labels":[],"label_agreement":null},{"id":"W4415974165","doi":"10.1016/j.procs.2025.09.440","title":"Towards automatic extraction of UML class diagrams: Creation of an annotated dataset for training deep models","year":2025,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal Military College of Canada","funders":"","keywords":"Unified Modeling Language; Class diagram; Automation; Applications of UML; Schema (genetic algorithms); Software; Structuring","score_opus":0.045972145385984264,"score_gpt":0.3260006784343023,"score_spread":0.28002853304831804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415974165","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23441379,0.0037759633,0.43713436,0.0020432072,0.0010265996,0.0016798856,0.25715247,0.042554148,0.020219686],"genre_scores_gemma":[0.15067239,0.0008123267,0.339808,0.00036520447,0.000092902534,0.002941581,0.49775928,0.0018021491,0.005746122],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99817026,0.0005560771,0.00019920658,0.00066907227,0.0002916748,0.000113696886],"domain_scores_gemma":[0.9931764,0.0038636057,0.0003503286,0.0009869756,0.0014054774,0.00021722629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023560566,0.0019595341,0.00071685994,0.00556691,0.00094319007,0.001439278,0.0016816924,0.00201662,0.005779827],"category_scores_gemma":[0.010151367,0.00065057556,0.0011404635,0.002862975,0.0007390537,0.0028187635,0.002014392,0.0024642914,0.005557678],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005919612,0.0012580493,0.016564969,0.003979321,0.00020965628,0.002389447,0.0025338335,0.028505372,0.05183275,0.011893132,0.2749335,0.6053081],"study_design_scores_gemma":[0.00035831658,0.00041659107,0.029210176,0.0011144185,0.0002973366,0.0016336117,0.0033208232,0.42000046,0.079015404,0.0176795,0.446705,0.00024836272],"about_ca_topic_score_codex":0.008403698,"about_ca_topic_score_gemma":0.015977561,"teacher_disagreement_score":0.008403698,"about_ca_system_score_codex":0.0014445488,"about_ca_system_score_gemma":0.0023665628,"threshold_uncertainty_score":0.019335449},"labels":[],"label_agreement":null},{"id":"W4415981209","doi":"10.48550/arxiv.2505.15962","title":"Pre-training Limited Memory Language Models with Internal and External Knowledge","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Division of Materials Research; Materials Research Science and Engineering Center, Harvard University; National Science Foundation; Advanced Research Projects Agency; Natural Sciences and Engineering Research Council of Canada; Cornell Center for Materials Research; Defense Advanced Research Projects Agency","keywords":"Memorization; Encoding (memory); Class (philosophy); Language model; Memory model; Knowledge acquisition; Knowledge representation and reasoning","score_opus":0.05873123808312963,"score_gpt":0.29150469282117875,"score_spread":0.2327734547380491,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415981209","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21516275,0.002126445,0.7458196,0.0022409307,0.0003652077,0.00021880292,0.0017769604,0.022611393,0.009677839],"genre_scores_gemma":[0.7977693,0.0006244036,0.1886161,0.0007911628,0.00015059295,0.00034175673,0.0037213033,0.000750226,0.007235228],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991844,0.00022896522,0.00006352195,0.00029662362,0.0001434761,0.000082911625],"domain_scores_gemma":[0.9941533,0.003639436,0.00022194034,0.0014217701,0.00041058965,0.0001528646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016407422,0.0013713039,0.00095832336,0.0006015079,0.00043274282,0.0020805567,0.0034699272,0.0015964038,0.004885897],"category_scores_gemma":[0.012949038,0.00080280926,0.0010428752,0.0007247329,0.00084084255,0.0071893483,0.0020736144,0.0037734124,0.0029555468],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008728449,0.00057320466,0.005595439,0.0006286498,0.00033177104,0.0004908304,0.00051520206,0.39067858,0.018626243,0.015060628,0.020038603,0.54658794],"study_design_scores_gemma":[0.000038585324,0.000094979885,0.00025914807,0.00003176289,0.000046497,0.000090636146,0.000055283577,0.9730334,0.0077024815,0.016148875,0.0024783672,0.000019985071],"about_ca_topic_score_codex":0.0041831774,"about_ca_topic_score_gemma":0.009360364,"teacher_disagreement_score":0.004885897,"about_ca_system_score_codex":0.0010847998,"about_ca_system_score_gemma":0.0015563101,"threshold_uncertainty_score":0.016344905},"labels":[],"label_agreement":null},{"id":"W4416016498","doi":"10.1145/3746252.3761322","title":"Enhancing and Assessing Instruction-Following with Fine-Grained Instruction Variants","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Robustness (evolution); Construct (python library); Benchmark (surveying); Training set; Focus (optics); Software deployment; Sensitivity (control systems)","score_opus":0.011822379688102767,"score_gpt":0.2519153502437036,"score_spread":0.24009297055560083,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416016498","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5192395,0.005192889,0.3686757,0.0012822964,0.0006361645,0.00048272472,0.013765434,0.08405538,0.006669891],"genre_scores_gemma":[0.73033875,0.0006246116,0.22261102,0.00090189505,0.00010815261,0.0005054772,0.03739438,0.003425497,0.0040903403],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983204,0.00044190546,0.00012008484,0.00075037,0.00024167888,0.00012556635],"domain_scores_gemma":[0.99616987,0.0020603896,0.00020084414,0.0009959592,0.00042664073,0.00014631759],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001812707,0.0021426587,0.0009168049,0.0013349481,0.0006064365,0.0015431327,0.0024705082,0.001340796,0.0023205876],"category_scores_gemma":[0.012216612,0.0006689372,0.0012420955,0.0012457127,0.0009827993,0.0035042095,0.0019258132,0.0028854252,0.0023104665],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014364505,0.0007520833,0.047158536,0.0014455783,0.0005504815,0.00044291024,0.0010568646,0.25163576,0.044371966,0.004369725,0.04596348,0.60081613],"study_design_scores_gemma":[0.00013075887,0.00046312134,0.004988758,0.00007246094,0.00015711068,0.00020607622,0.0004110367,0.9370172,0.031873338,0.009012409,0.015576859,0.00009098646],"about_ca_topic_score_codex":0.008258961,"about_ca_topic_score_gemma":0.021673767,"teacher_disagreement_score":0.008258961,"about_ca_system_score_codex":0.0009686688,"about_ca_system_score_gemma":0.0018191369,"threshold_uncertainty_score":0.016421795},"labels":[],"label_agreement":null},{"id":"W4416016556","doi":"10.1145/3746252.3761587","title":"<scp>FinSage:</scp> A Multi-aspect RAG System for Financial Filings Question Answering","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal; Mila - Quebec Artificial Intelligence Institute; Ontario Brain Institute","funders":"","keywords":"Automatic summarization; Pipeline (software); Metadata; Question answering; Precision and recall; XBRL; Order (exchange); Baseline (sea); Semantic heterogeneity","score_opus":0.0198375242962318,"score_gpt":0.26681399796033856,"score_spread":0.24697647366410674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416016556","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006316997,0.00042474648,0.3448647,0.0022251133,0.00036849632,0.0007187989,0.026822,0.6047517,0.013507448],"genre_scores_gemma":[0.13565885,0.00039846558,0.63140047,0.0038022595,0.00042421537,0.0013386877,0.16507716,0.026313432,0.035586458],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99867284,0.00028273833,0.00012091952,0.00037804202,0.00043902497,0.00010635875],"domain_scores_gemma":[0.9974113,0.0010116239,0.00009820525,0.0006121788,0.00067802996,0.00018869882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024118517,0.0018930789,0.0008199817,0.0016264996,0.0006930804,0.0020673403,0.0029292058,0.002318682,0.053149115],"category_scores_gemma":[0.009353191,0.0007166758,0.0014621504,0.0008773433,0.0009125968,0.004906879,0.0037849233,0.0020009037,0.029899461],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037895865,0.00017413762,0.001313131,0.000753346,0.00010331037,0.00063794345,0.00066226517,0.0051291003,0.020088064,0.008854882,0.78053683,0.18136798],"study_design_scores_gemma":[0.00028958134,0.00027990495,0.0025970014,0.00013641751,0.00007126622,0.0009859465,0.0004146444,0.34665567,0.03587945,0.027410131,0.5850941,0.00018586271],"about_ca_topic_score_codex":0.009714048,"about_ca_topic_score_gemma":0.014167963,"teacher_disagreement_score":0.053149115,"about_ca_system_score_codex":0.0011154962,"about_ca_system_score_gemma":0.0014465393,"threshold_uncertainty_score":0.17780149},"labels":[],"label_agreement":null},{"id":"W4416017155","doi":"10.1145/3746252.3760960","title":"Study on LLMs for Promptagator-Style Dense Retriever Training","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Work (physics); Code (set theory); Training (meteorology); Risk assessment; Government (linguistics)","score_opus":0.09112275622836809,"score_gpt":0.32568286687121645,"score_spread":0.23456011064284837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416017155","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2878713,0.0058264015,0.5917907,0.002454331,0.000774943,0.0009878116,0.005643936,0.08583393,0.01881663],"genre_scores_gemma":[0.63907677,0.00087895175,0.32964307,0.0016913429,0.00022626638,0.00083442574,0.01527115,0.0042287107,0.008149405],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99768376,0.0010881247,0.00013606656,0.0005935992,0.00033157566,0.00016680131],"domain_scores_gemma":[0.99026793,0.00723428,0.00015659371,0.001382717,0.00070471974,0.0002537042],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004660524,0.0015752609,0.0010814767,0.0008054233,0.00068038737,0.0016955909,0.002250608,0.0015939607,0.007519622],"category_scores_gemma":[0.023186339,0.00059382786,0.00091223646,0.00093568757,0.00076010247,0.003937832,0.0019235343,0.0032344162,0.0050840685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001978515,0.0013019558,0.007358773,0.0019277347,0.00045474453,0.0006936311,0.0011323105,0.23951322,0.03309617,0.013092379,0.08024541,0.6192051],"study_design_scores_gemma":[0.00030869746,0.0005621959,0.001639859,0.00007080932,0.00007845303,0.00033954324,0.0004030388,0.95197994,0.018606456,0.009185368,0.016764963,0.00006077414],"about_ca_topic_score_codex":0.008154227,"about_ca_topic_score_gemma":0.012943676,"teacher_disagreement_score":0.008154227,"about_ca_system_score_codex":0.0013327644,"about_ca_system_score_gemma":0.001553074,"threshold_uncertainty_score":0.025155663},"labels":[],"label_agreement":null},{"id":"W4416017238","doi":"10.1145/3746252.3760922","title":"LLM-as-a-Judge in Entity Retrieval: Assessing Explicit and Implicit Relevance","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Toronto Metropolitan University","funders":"","keywords":"Relevance (law); Context (archaeology); Replicate; Benchmark (surveying); Reliability (semiconductor)","score_opus":0.023802537527508855,"score_gpt":0.31273081034597494,"score_spread":0.2889282728184661,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416017238","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60237104,0.008370657,0.28630805,0.0026870084,0.0010150407,0.0012545002,0.01909197,0.057695378,0.021206308],"genre_scores_gemma":[0.80923784,0.00042599015,0.16025923,0.00070860604,0.00022353849,0.00046642992,0.022941863,0.0015318639,0.0042045508],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98802316,0.006680356,0.00080081436,0.0024777323,0.0014902051,0.0005278484],"domain_scores_gemma":[0.95432943,0.032254677,0.001996611,0.0063816006,0.0035977613,0.0014398766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017069725,0.002047289,0.0015341253,0.0041667866,0.0012253433,0.0038500873,0.0022956128,0.0032860474,0.0041502067],"category_scores_gemma":[0.08118684,0.0006139582,0.0011463014,0.002315905,0.0010976817,0.0067369635,0.0041729817,0.0028798976,0.0044817827],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008020369,0.0018701908,0.20270202,0.00384004,0.0014491463,0.00089819124,0.005524239,0.0775982,0.027174527,0.01384271,0.1199189,0.5371615],"study_design_scores_gemma":[0.00034085297,0.0007049069,0.033468205,0.000243865,0.0001923614,0.00042544075,0.0012583131,0.90232396,0.018655663,0.016954498,0.025164602,0.00026726964],"about_ca_topic_score_codex":0.01133637,"about_ca_topic_score_gemma":0.021212589,"teacher_disagreement_score":0.017069725,"about_ca_system_score_codex":0.0013742717,"about_ca_system_score_gemma":0.0017634706,"threshold_uncertainty_score":0.09027445},"labels":[],"label_agreement":null},{"id":"W4416017552","doi":"10.1145/3746252.3761355","title":"Evolving Graph-Based Context Modeling for Multi-Turn Conversational Retrieval-Augmented Generation","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"National Science and Technology Major Project; National Natural Science Foundation of China","keywords":"Leverage (statistics); Rewriting; Knowledge graph; Context (archaeology); Graph; Representation (politics); Key (lock); Knowledge representation and reasoning","score_opus":0.08851693343086302,"score_gpt":0.3039273201317528,"score_spread":0.21541038670088974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416017552","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01720968,0.0011774339,0.97422296,0.00023051331,0.00006376691,0.00013878597,0.00046159216,0.0048344075,0.0016608201],"genre_scores_gemma":[0.54653394,0.00087431865,0.44305062,0.00040610548,0.00013432972,0.00038386206,0.0031466498,0.00073022,0.0047398605],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914503,0.0003282149,0.000034968478,0.00028822405,0.00013810703,0.000065347536],"domain_scores_gemma":[0.9990563,0.00057519815,0.00005673275,0.00012661608,0.00014000655,0.000045138477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008725076,0.0011288272,0.0008022518,0.0011983732,0.0005452995,0.00080102234,0.0017204586,0.0009993188,0.0021834571],"category_scores_gemma":[0.0032114692,0.0004502003,0.0010734111,0.0008223017,0.00049161667,0.0014492467,0.0012649664,0.0012589325,0.0010277815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046769337,0.00019788074,0.0018437945,0.00043427796,0.00019335096,0.0004399072,0.00094366126,0.5152708,0.025585862,0.01634898,0.011725254,0.42654848],"study_design_scores_gemma":[0.00001810906,0.00003498395,0.0002176548,0.000007603327,0.000029126168,0.0000554077,0.000048092472,0.9859489,0.0024400346,0.0083690025,0.002816085,0.000015015273],"about_ca_topic_score_codex":0.01346526,"about_ca_topic_score_gemma":0.0220148,"teacher_disagreement_score":0.01346526,"about_ca_system_score_codex":0.000925567,"about_ca_system_score_gemma":0.0010095071,"threshold_uncertainty_score":0.02677375},"labels":[],"label_agreement":null},{"id":"W4416017611","doi":"10.1145/3746252.3761593","title":"ProActLLM: Proactive Conversational Information Seeking with Large Language Models","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Information seeking; Cognition; Language model; Key (lock); Action (physics); Information access; Cognitive model; Proactivity","score_opus":0.010335465320123442,"score_gpt":0.22697544565583652,"score_spread":0.21663998033571308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416017611","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012763913,0.00077402574,0.9572544,0.0020271307,0.00035012834,0.00045540568,0.00080672104,0.01680108,0.008767244],"genre_scores_gemma":[0.16141255,0.00069183816,0.8112678,0.00091556384,0.00027963452,0.0013043518,0.0037272177,0.00290539,0.017495591],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937528,0.0041670585,0.0002481676,0.00081328937,0.00083525525,0.00018340067],"domain_scores_gemma":[0.98695254,0.009648909,0.0002988871,0.0019069262,0.0006311205,0.0005616697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008414199,0.001222468,0.00097369176,0.0010199298,0.0017311844,0.006085627,0.0034696946,0.0031027675,0.01375744],"category_scores_gemma":[0.02025049,0.0012603573,0.0015403345,0.0007237562,0.0012132551,0.010548045,0.0073681516,0.0039013342,0.005344095],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022563804,0.0010287199,0.0028617387,0.002031801,0.00050202693,0.0013149738,0.016263247,0.0312108,0.047366653,0.19272965,0.13362063,0.5688134],"study_design_scores_gemma":[0.00031549128,0.00040397386,0.00069363415,0.0001981234,0.00012873572,0.0005281487,0.0019441767,0.6341501,0.018494349,0.11530849,0.227635,0.0001997952],"about_ca_topic_score_codex":0.0023651044,"about_ca_topic_score_gemma":0.0041042147,"teacher_disagreement_score":0.01375744,"about_ca_system_score_codex":0.0010135556,"about_ca_system_score_gemma":0.0017170674,"threshold_uncertainty_score":0.04602319},"labels":[],"label_agreement":null},{"id":"W4416033588","doi":"10.18653/v1/2025.newsum-main.4","title":"Beyond Paraphrasing: Analyzing Summarization Abstractiveness and Reasoning","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institut de Valorisation des Données; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Automatic summarization; Feature (linguistics); Key (lock); Term (time); Context (archaeology)","score_opus":0.013495264458107389,"score_gpt":0.269114633124226,"score_spread":0.2556193686661186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416033588","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49962494,0.0040382906,0.46385938,0.0019131574,0.00011278047,0.0009267625,0.0046578865,0.004041812,0.020824911],"genre_scores_gemma":[0.90617883,0.0006559412,0.08673936,0.0001747633,0.00005459026,0.00017457092,0.003919746,0.00027869205,0.0018235103],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9939446,0.0034680755,0.00055601425,0.00076006475,0.0010974533,0.00017376707],"domain_scores_gemma":[0.90190244,0.07234467,0.010186215,0.008138025,0.0066973246,0.0007313491],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008921248,0.0006334314,0.0006058935,0.0046768966,0.0007517228,0.0041963933,0.0011200145,0.0011195985,0.0036917198],"category_scores_gemma":[0.09356151,0.00032008888,0.0007571955,0.00413158,0.0010649982,0.007988357,0.0015113365,0.0014804623,0.00089903944],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016034864,0.0005765396,0.11408327,0.0048923534,0.0009997905,0.0010829124,0.057299163,0.05855497,0.04951565,0.081559524,0.017245742,0.61258656],"study_design_scores_gemma":[0.00018809507,0.000979519,0.10460571,0.0009913507,0.0008998598,0.0012489324,0.015910152,0.6167969,0.035878558,0.1557343,0.066420145,0.00034647883],"about_ca_topic_score_codex":0.0035700204,"about_ca_topic_score_gemma":0.003097729,"teacher_disagreement_score":0.008921248,"about_ca_system_score_codex":0.0011724103,"about_ca_system_score_gemma":0.00069471454,"threshold_uncertainty_score":0.047180653},"labels":[],"label_agreement":null},{"id":"W4416034334","doi":"10.18653/v1/2025.findings-emnlp.846","title":"Not Lost After All: How Cross-Encoder Attribution Challenges Position Bias Assumptions in LLM Summarization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Killam Trusts; Canadian Institute for Advanced Research","keywords":"Position (finance); Automatic summarization; Attribution; Position paper; Term (time)","score_opus":0.08732728277364163,"score_gpt":0.3179720882859166,"score_spread":0.23064480551227498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416034334","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049831815,0.002178644,0.92467344,0.007933036,0.0011604578,0.00013709332,0.0015381987,0.005119001,0.007428396],"genre_scores_gemma":[0.8214539,0.0010485023,0.16066782,0.001963094,0.0014808698,0.00020014559,0.0035134065,0.0018403603,0.007831806],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9901557,0.005622092,0.00060165947,0.0019945202,0.001128506,0.0004974699],"domain_scores_gemma":[0.94375366,0.04034318,0.0019507315,0.008503796,0.0047359373,0.00071273046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01338912,0.0011361651,0.0014287034,0.0016686592,0.0016589882,0.0053517087,0.002240558,0.0032730382,0.0059707677],"category_scores_gemma":[0.096821405,0.0009417632,0.00083594734,0.002162638,0.0014734287,0.010551386,0.004165392,0.0055724108,0.0043844883],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001639825,0.00023383812,0.011287682,0.00088690425,0.00052506255,0.0005455009,0.002858616,0.039662328,0.012379271,0.068869516,0.04830765,0.81280386],"study_design_scores_gemma":[0.00013204257,0.00026402998,0.004199613,0.0004078162,0.0003831842,0.00053344603,0.00117822,0.72170126,0.02165959,0.22813846,0.021262195,0.00014009514],"about_ca_topic_score_codex":0.0043154377,"about_ca_topic_score_gemma":0.004755464,"teacher_disagreement_score":0.01338912,"about_ca_system_score_codex":0.0011216566,"about_ca_system_score_gemma":0.0022327823,"threshold_uncertainty_score":0.070809305},"labels":[],"label_agreement":null},{"id":"W4416034841","doi":"10.18653/v1/2025.findings-emnlp.662","title":"Topic-Guided Reinforcement Learning with LLMs for Enhancing Multi-Document Summarization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Reinforcement learning; Reinforcement; Action (physics)","score_opus":0.024920055596030465,"score_gpt":0.28843534458846215,"score_spread":0.2635152889924317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416034841","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027066588,0.0010156521,0.9678717,0.00030960253,0.00007393491,0.0000674978,0.0001288795,0.0026118045,0.0008542703],"genre_scores_gemma":[0.700952,0.00055717764,0.2921865,0.00040585108,0.0003077134,0.00025805485,0.00087620533,0.0004519899,0.004004411],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99916375,0.00032479834,0.00005273936,0.00023260822,0.00015876426,0.00006731162],"domain_scores_gemma":[0.997621,0.0015312411,0.00024623208,0.00017812962,0.00032064525,0.00010277213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017649756,0.00095714664,0.0013873316,0.0008773245,0.00039914344,0.0008776685,0.0012276564,0.0009892267,0.0014621107],"category_scores_gemma":[0.00612138,0.00034893295,0.0005131324,0.0009361432,0.00045801222,0.0018862392,0.0011091991,0.0014158586,0.00072062825],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003842869,0.00028990625,0.0019004233,0.00035065212,0.00013959914,0.00016763533,0.00053551147,0.52158135,0.020871475,0.0070909956,0.0069937143,0.4396944],"study_design_scores_gemma":[0.000023201434,0.00006406048,0.0001396205,0.000008281077,0.000017928884,0.000015202237,0.000019281608,0.99262744,0.0023401633,0.0037847338,0.0009516053,0.000008573821],"about_ca_topic_score_codex":0.0025701441,"about_ca_topic_score_gemma":0.0043386663,"teacher_disagreement_score":0.0025701441,"about_ca_system_score_codex":0.00079874624,"about_ca_system_score_gemma":0.0010106402,"threshold_uncertainty_score":0.009334207},"labels":[],"label_agreement":null},{"id":"W4416035393","doi":"10.18653/v1/2025.emnlp-main.1506","title":"Machine-generated text detection prevents language model collapse","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Engineering and Physical Sciences Research Council","keywords":"Language model; Natural language; Language identification; Noise (video); Field (mathematics)","score_opus":0.0121339640939598,"score_gpt":0.25614073424940154,"score_spread":0.24400677015544173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416035393","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36017027,0.001020002,0.61684245,0.0008449104,0.0002851067,0.00028374366,0.00060089736,0.017411118,0.002541429],"genre_scores_gemma":[0.8134007,0.00017332069,0.17976137,0.0005425174,0.00010871378,0.00021062861,0.0024918604,0.0009808157,0.002330065],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970298,0.0011674045,0.00017445393,0.000868932,0.0005378364,0.00022156487],"domain_scores_gemma":[0.9806486,0.012242264,0.0010474634,0.0032348996,0.0023583104,0.0004685529],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060591586,0.001847034,0.0014294792,0.00143859,0.0009419718,0.0015455444,0.001983605,0.0021409418,0.0011441573],"category_scores_gemma":[0.03934618,0.0006597677,0.0010475682,0.00067721744,0.0014293172,0.0029700736,0.0024218676,0.0027140083,0.0017168158],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013296842,0.0008209496,0.021035535,0.00058573124,0.00039376447,0.0007593759,0.0011322786,0.4170237,0.062451977,0.0052644396,0.012963842,0.47623867],"study_design_scores_gemma":[0.000036739762,0.00018939975,0.0013257646,0.000027716815,0.00003184766,0.00018084858,0.000077139244,0.95946264,0.035145458,0.0019411735,0.0015537542,0.000027501173],"about_ca_topic_score_codex":0.004334304,"about_ca_topic_score_gemma":0.005318798,"teacher_disagreement_score":0.0060591586,"about_ca_system_score_codex":0.0010099054,"about_ca_system_score_gemma":0.0019978553,"threshold_uncertainty_score":0.032044232},"labels":[],"label_agreement":null},{"id":"W4416035419","doi":"10.18653/v1/2025.emnlp-main.1605","title":"Are Stereotypes Leading LLMs’ Zero-Shot Stance Detection ?","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Agence Nationale de la Recherche","keywords":"Perspective (graphical); Noise (video); Perception; Action (physics)","score_opus":0.044488391414404,"score_gpt":0.28302415477266596,"score_spread":0.23853576335826196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416035419","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79219466,0.003358053,0.1742694,0.0039172573,0.000849974,0.00021620793,0.0045592887,0.0067454665,0.01388979],"genre_scores_gemma":[0.95714694,0.00026267034,0.033403136,0.0013018722,0.00015676611,0.00010259269,0.0046428447,0.00048570413,0.0024973906],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99546623,0.0023784041,0.00023588323,0.0010506002,0.00055441837,0.00031437405],"domain_scores_gemma":[0.98473144,0.010421865,0.00094438187,0.0022288857,0.0011762525,0.00049718533],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068682064,0.0010841917,0.0008676032,0.0010269891,0.0010137786,0.0022482045,0.001072318,0.0013392468,0.001968982],"category_scores_gemma":[0.031182788,0.0005126787,0.0006408799,0.00057265384,0.0009667694,0.0032872455,0.0019096131,0.0021801372,0.0029790862],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004205468,0.00077659945,0.24175514,0.0014262726,0.001032144,0.0014307287,0.008063391,0.025169596,0.0890079,0.009538696,0.053887106,0.56370705],"study_design_scores_gemma":[0.00030722903,0.0010231611,0.10493941,0.00048428058,0.00036474623,0.0020336453,0.00779596,0.6852639,0.0839897,0.07100126,0.04251971,0.00027709836],"about_ca_topic_score_codex":0.004511107,"about_ca_topic_score_gemma":0.0123174,"teacher_disagreement_score":0.0068682064,"about_ca_system_score_codex":0.0008183838,"about_ca_system_score_gemma":0.00075459655,"threshold_uncertainty_score":0.03632301},"labels":[],"label_agreement":null},{"id":"W4416035446","doi":"10.18653/v1/2025.emnlp-main.1398","title":"Trustworthy Medical Question Answering: An Evaluation-Centric Survey","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Science Foundation of Shandong Province; Natural Sciences and Engineering Research Council of Canada; National Natural Science Foundation of China; Canadian Institute for Advanced Research","keywords":"Trustworthiness; Natural (archaeology); Empirical research; Questions and answers; Natural language","score_opus":0.055839135060534985,"score_gpt":0.35406033832430484,"score_spread":0.29822120326376983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416035446","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15466987,0.6130073,0.08433972,0.091836564,0.0014590868,0.0011428713,0.007221902,0.0020595437,0.044263147],"genre_scores_gemma":[0.8550937,0.09449416,0.027351514,0.009632669,0.0020867065,0.00052100816,0.008508604,0.00049280334,0.0018187277],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.91828996,0.043627758,0.006665472,0.0046701343,0.025201354,0.0015453653],"domain_scores_gemma":[0.4163363,0.5109869,0.0178809,0.015764628,0.035521038,0.0035101396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.105293006,0.00070985843,0.0016446295,0.0102544855,0.0011610361,0.0039010695,0.0030183152,0.0029446152,0.00331608],"category_scores_gemma":[0.32645908,0.00064825645,0.0009799952,0.006909647,0.002666224,0.009119486,0.0038539057,0.002078989,0.001243114],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010837866,0.0003219853,0.11740821,0.007405577,0.0006843519,0.00019045931,0.0039320714,0.0021328367,0.0014360365,0.009884589,0.063638575,0.7918815],"study_design_scores_gemma":[0.00064451405,0.0027910734,0.28516313,0.024595385,0.00338871,0.0063808965,0.013938183,0.07008335,0.011059573,0.06730167,0.5142449,0.00040857817],"about_ca_topic_score_codex":0.0035182973,"about_ca_topic_score_gemma":0.003982743,"teacher_disagreement_score":0.105293006,"about_ca_system_score_codex":0.0029163074,"about_ca_system_score_gemma":0.0034245702,"threshold_uncertainty_score":0.5568493},"labels":[],"label_agreement":null},{"id":"W4416036138","doi":"10.18653/v1/2025.emnlp-main.1075","title":"Improving Context Fidelity via Native Retrieval-Augmented Reasoning","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada","keywords":"Context (archaeology); Natural language; Fidelity; Natural (archaeology); Natural language understanding","score_opus":0.01701550256221682,"score_gpt":0.27022452701014765,"score_spread":0.2532090244479308,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416036138","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08223024,0.006525199,0.88192546,0.0006027843,0.0006383535,0.00027242495,0.0009463243,0.015556341,0.011302799],"genre_scores_gemma":[0.65320086,0.0014715974,0.33534855,0.00057163805,0.00031832064,0.00012368472,0.00204763,0.000751908,0.0061658947],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975067,0.0005784072,0.00019414796,0.0006971765,0.00077640446,0.00024722685],"domain_scores_gemma":[0.9955604,0.0015960485,0.00015379782,0.0019144146,0.00065301516,0.00012236307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019603353,0.0016037173,0.00240186,0.0017380269,0.0009826476,0.0025356866,0.0026257494,0.0018570618,0.0076413127],"category_scores_gemma":[0.012269021,0.00064246386,0.0013294327,0.0014191341,0.00071520824,0.0069169602,0.0045742462,0.0019443858,0.0028617666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013772323,0.00078662275,0.0028243815,0.00078525813,0.00026588968,0.0004101006,0.0006343838,0.045526646,0.04123369,0.012734559,0.023449836,0.86997133],"study_design_scores_gemma":[0.00032955012,0.00036613294,0.0015670824,0.00009145308,0.00049022347,0.0005557906,0.00043249037,0.9220186,0.029439922,0.030289548,0.014320173,0.00009907542],"about_ca_topic_score_codex":0.0066662882,"about_ca_topic_score_gemma":0.010060501,"teacher_disagreement_score":0.0076413127,"about_ca_system_score_codex":0.00057922903,"about_ca_system_score_gemma":0.0015536286,"threshold_uncertainty_score":0.025562763},"labels":[],"label_agreement":null},{"id":"W4416036180","doi":"10.18653/v1/2025.emnlp-main.830","title":"Not-Just-Scaling Laws: Towards a Better Understanding of the Downstream Impact of Language Model Design Decisions","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Institute of Advanced Industrial Science and Technology; National Science Foundation","keywords":"Downstream (manufacturing); Natural language; Language model; Natural (archaeology); Empirical research; Natural language generation","score_opus":0.12442220833539444,"score_gpt":0.34005445658991074,"score_spread":0.2156322482545163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416036180","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04662633,0.0034473555,0.88181514,0.021438964,0.000615136,0.00015728164,0.00066015706,0.0016323017,0.04360729],"genre_scores_gemma":[0.771352,0.002813686,0.21056308,0.0033693465,0.0006478505,0.000411442,0.0007998268,0.0024585587,0.0075841257],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99339676,0.0037623006,0.00028455467,0.0009803607,0.0012688959,0.00030714308],"domain_scores_gemma":[0.88704395,0.08802734,0.0048955777,0.013610648,0.005097416,0.0013250852],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011541389,0.0011128851,0.0014554345,0.0012453508,0.0014065111,0.006591367,0.0025389702,0.0023795995,0.01643264],"category_scores_gemma":[0.14354348,0.0013004843,0.0010237445,0.0017068402,0.0035861232,0.027453028,0.0038010385,0.007346522,0.0022976615],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017573965,0.00022571057,0.005312537,0.0003133291,0.00011218927,0.00020710035,0.0011894078,0.03751858,0.0016407433,0.8461469,0.019991418,0.08716634],"study_design_scores_gemma":[0.000023110539,0.000025131942,0.00056357175,0.000050904975,0.000025360257,0.000040734,0.00017199774,0.12147052,0.0009213039,0.87215453,0.004526494,0.000026319794],"about_ca_topic_score_codex":0.004449144,"about_ca_topic_score_gemma":0.004215486,"teacher_disagreement_score":0.01643264,"about_ca_system_score_codex":0.0018235833,"about_ca_system_score_gemma":0.0022805799,"threshold_uncertainty_score":0.06103736},"labels":[],"label_agreement":null},{"id":"W4416036370","doi":"10.18653/v1/2025.emnlp-main.590","title":"CEMTM: Contextual Embedding-based Multimodal Topic Modeling","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Topic model; Feature (linguistics); Field (mathematics); Key (lock); Context (archaeology)","score_opus":0.030170760914198732,"score_gpt":0.2980552092464189,"score_spread":0.26788444833222014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416036370","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013447104,0.0017866419,0.977178,0.00027473166,0.00011944212,0.000107344604,0.0015636187,0.0038244447,0.0016987623],"genre_scores_gemma":[0.440617,0.0027184917,0.53659856,0.0006049244,0.00054964947,0.0007802493,0.009240293,0.0011496221,0.0077411947],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995277,0.00012705814,0.000027855662,0.000175821,0.00008529378,0.000056210505],"domain_scores_gemma":[0.99937314,0.00029948194,0.000056528796,0.0001058835,0.0001315879,0.000033360677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008323546,0.001220983,0.00076020753,0.0013369686,0.00031973506,0.0010913064,0.0013363474,0.0009217405,0.0034607828],"category_scores_gemma":[0.0033662894,0.00035263656,0.0012219783,0.0014738062,0.00038987564,0.002240883,0.0015576355,0.0015070302,0.0019004145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005194713,0.00018512696,0.0025677185,0.00073363393,0.00028497804,0.00036036968,0.0007240897,0.14719944,0.03148256,0.020922389,0.027655361,0.7673648],"study_design_scores_gemma":[0.00003095455,0.000100727826,0.0008293895,0.000058509228,0.00008478114,0.00021496392,0.00011437854,0.95777726,0.008497976,0.020874059,0.01137625,0.00004070486],"about_ca_topic_score_codex":0.004740887,"about_ca_topic_score_gemma":0.0068660537,"teacher_disagreement_score":0.004740887,"about_ca_system_score_codex":0.00063472684,"about_ca_system_score_gemma":0.0007030322,"threshold_uncertainty_score":0.011577487},"labels":[],"label_agreement":null},{"id":"W4416036669","doi":"10.18653/v1/2025.emnlp-main.368","title":"LEO-MINI: An Efficient Multimodal Large Language Model using Conditional Token Reduction and Mixture of Multi-Modal Experts","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Government of Canada","keywords":"Reduction (mathematics); Language model; Security token; Data modeling; Natural language","score_opus":0.027423990619345126,"score_gpt":0.31452642516473067,"score_spread":0.28710243454538553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416036669","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071244496,0.0003244411,0.9829838,0.00024271515,0.000058403264,0.000103522696,0.0005232801,0.0076341447,0.0010051257],"genre_scores_gemma":[0.26176068,0.0004364958,0.71967554,0.0010594279,0.00013992684,0.00071717426,0.004768587,0.0012801021,0.010162106],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99946445,0.00013989421,0.000029849349,0.00018560584,0.000112820126,0.00006732353],"domain_scores_gemma":[0.9991191,0.000493484,0.00004699167,0.00012157681,0.00015714898,0.00006169156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011741834,0.0012197752,0.0010365364,0.00065864314,0.00043838497,0.0011232803,0.0028346963,0.0011821458,0.004149754],"category_scores_gemma":[0.0032740158,0.0007559381,0.0015605422,0.00059436884,0.0005078091,0.002648385,0.002135461,0.0026347071,0.00218143],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00071450754,0.00034315564,0.0013106354,0.00040974992,0.0003363227,0.00030736448,0.00028377745,0.35802802,0.022309711,0.018703315,0.029212443,0.56804097],"study_design_scores_gemma":[0.000016101318,0.000030562045,0.00008645724,0.0000064001806,0.000017280438,0.000032471104,0.000016019285,0.9908982,0.0024589226,0.0048207,0.0016028049,0.000014117676],"about_ca_topic_score_codex":0.008075685,"about_ca_topic_score_gemma":0.016569454,"teacher_disagreement_score":0.008075685,"about_ca_system_score_codex":0.0009738197,"about_ca_system_score_gemma":0.0016794757,"threshold_uncertainty_score":0.016057372},"labels":[],"label_agreement":null},{"id":"W4416036956","doi":"10.18653/v1/2025.emnlp-main.285","title":"Tree-of-Quote Prompting Improves Factuality and Attribution in Multi-Hop and Medical Reasoning","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institutes of Health Research; Deutsche Forschungsgemeinschaft; University of Oxford; Johns Hopkins University","keywords":"Attribution; Natural (archaeology); Natural language; Empirical research; Empirical evidence","score_opus":0.03687982815593502,"score_gpt":0.31707948615818476,"score_spread":0.2801996580022498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416036956","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3066036,0.005761757,0.6271674,0.008437752,0.0018689637,0.00072089134,0.0034110725,0.02936761,0.016661014],"genre_scores_gemma":[0.7822768,0.0007716443,0.20951705,0.00060631667,0.00026265823,0.00007439002,0.0028357133,0.0004759361,0.0031793488],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99588734,0.0019776148,0.00036913264,0.0009213815,0.00067532196,0.00016926673],"domain_scores_gemma":[0.95823777,0.03197791,0.002035452,0.004216541,0.002383114,0.001149206],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076483167,0.00079659326,0.0009993379,0.001880234,0.0010419689,0.002155256,0.0016482838,0.0019076103,0.01099647],"category_scores_gemma":[0.054885026,0.00048641188,0.00083639944,0.001317641,0.0006321233,0.008304112,0.0034284522,0.0024223756,0.0023721692],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028378624,0.0011731453,0.016452994,0.0007058345,0.00018013958,0.00031875542,0.0015198909,0.019134764,0.0057489933,0.012571604,0.03002147,0.9093346],"study_design_scores_gemma":[0.0008068041,0.0010228318,0.015443408,0.00047083627,0.00069540617,0.0005982231,0.0018472649,0.72053427,0.028175665,0.20045651,0.029746324,0.00020247462],"about_ca_topic_score_codex":0.0035363752,"about_ca_topic_score_gemma":0.005981509,"teacher_disagreement_score":0.01099647,"about_ca_system_score_codex":0.00081057963,"about_ca_system_score_gemma":0.0021661187,"threshold_uncertainty_score":0.040448666},"labels":[],"label_agreement":null},{"id":"W4416037923","doi":"10.18653/v1/2025.codi-1.13","title":"Discourse Relation Recognition with Language Models Under Different Data Availability","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Relation (database); Language model; Natural language; Feature (linguistics); Language identification; Data modeling","score_opus":0.08687933868806245,"score_gpt":0.31176919319229723,"score_spread":0.22488985450423477,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416037923","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37791088,0.0022520232,0.59429073,0.002665231,0.0005758154,0.00024709245,0.0054522487,0.0124361925,0.0041698324],"genre_scores_gemma":[0.8894743,0.00046306854,0.100597896,0.00016830597,0.00021160101,0.00019166066,0.006262531,0.00047574556,0.0021548267],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99599934,0.0013943985,0.0003255295,0.0013364808,0.0005138131,0.00043050232],"domain_scores_gemma":[0.98647434,0.009269358,0.00032734164,0.0023902839,0.0012134389,0.00032508024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048447773,0.0011006005,0.0017116773,0.0015061041,0.0010066729,0.0032347313,0.0017140403,0.0015665912,0.0034893982],"category_scores_gemma":[0.021567836,0.00060242444,0.001520093,0.001661757,0.0006227793,0.0069004847,0.0019231699,0.0026910384,0.0023160481],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0065697264,0.0013964023,0.01402386,0.00091225,0.00069149205,0.00071201427,0.0015236908,0.17323199,0.047075372,0.010803199,0.01619007,0.72687],"study_design_scores_gemma":[0.00006941765,0.00008621231,0.0015081571,0.000021951197,0.0001416544,0.000078741854,0.0003240047,0.9726458,0.013186485,0.010159403,0.001733596,0.00004446671],"about_ca_topic_score_codex":0.009848559,"about_ca_topic_score_gemma":0.008698337,"teacher_disagreement_score":0.009848559,"about_ca_system_score_codex":0.0010144458,"about_ca_system_score_gemma":0.002187956,"threshold_uncertainty_score":0.025621891},"labels":[],"label_agreement":null},{"id":"W4416075695","doi":"10.31224/5781","title":"Retrieval-Augmented Generation","year":2025,"lang":"","type":"preprint","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Natural language generation; Key (lock); Generative grammar; Natural language; Context (archaeology); Language model; Information extraction; Architecture","score_opus":0.06732035046122521,"score_gpt":0.2972475834883907,"score_spread":0.2299272330271655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416075695","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012670689,0.0019687242,0.9158909,0.00063249696,0.00037811656,0.0006202894,0.0011681718,0.05174128,0.014929356],"genre_scores_gemma":[0.27513453,0.002001421,0.684676,0.0012398133,0.00034210068,0.00064526853,0.0060033896,0.0048036347,0.025153859],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980096,0.0006233742,0.00014156019,0.00056175445,0.00051102124,0.00015270904],"domain_scores_gemma":[0.9950864,0.0017181057,0.00020736211,0.0021790748,0.00066768425,0.00014138076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027859001,0.0015883566,0.0012922555,0.0014438124,0.00080152083,0.0029143502,0.00349666,0.0017209919,0.016415037],"category_scores_gemma":[0.008758585,0.00083026505,0.0016391409,0.001451717,0.0013365307,0.004801584,0.004427638,0.0018521838,0.008908238],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009569739,0.00046541123,0.0028774906,0.001330214,0.0003209696,0.0007995785,0.0012483316,0.038457483,0.029761989,0.059675243,0.057433248,0.806673],"study_design_scores_gemma":[0.00040052904,0.0006307447,0.0016121875,0.00024133234,0.000455509,0.0016508098,0.0003179189,0.59188664,0.06855685,0.12734677,0.20661905,0.00028167423],"about_ca_topic_score_codex":0.003204286,"about_ca_topic_score_gemma":0.002999726,"teacher_disagreement_score":0.016415037,"about_ca_system_score_codex":0.00087730354,"about_ca_system_score_gemma":0.0016858734,"threshold_uncertainty_score":0.05491376},"labels":[],"label_agreement":null},{"id":"W4416095464","doi":"10.48550/arxiv.2505.04666","title":"Fine-Tuning Large Language Models and Evaluating Retrieval Methods for Improved Question Answering on Building Codes","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Context (archaeology); Language model; Question answering; Code (set theory); Key (lock); Protocol (science); Adaptation (eye)","score_opus":0.10268088302087734,"score_gpt":0.4276200098373957,"score_spread":0.3249391268165183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416095464","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6088455,0.008546866,0.33264777,0.0020898764,0.00066460297,0.0022341092,0.0051427605,0.0324005,0.0074280826],"genre_scores_gemma":[0.67571515,0.0015187785,0.29625854,0.001062025,0.00024320949,0.0010673888,0.01965527,0.00088890083,0.0035908318],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99368733,0.003348576,0.0005482232,0.001421084,0.00070097984,0.00029377945],"domain_scores_gemma":[0.9790044,0.016961964,0.00052564766,0.0017373126,0.0014605648,0.0003101318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009238964,0.0022251182,0.0015575909,0.0023733207,0.00072256016,0.001894548,0.0029548204,0.0020693275,0.0025007068],"category_scores_gemma":[0.029025454,0.0005298647,0.0019733582,0.0013956178,0.00091998046,0.00470218,0.0017366393,0.0030330068,0.0018261786],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026002622,0.0025974454,0.009625354,0.0035295924,0.0007688402,0.00062198774,0.0016566382,0.22867492,0.05404207,0.0031282357,0.024757996,0.66799664],"study_design_scores_gemma":[0.0004128341,0.0009900897,0.0041251513,0.00010586736,0.00026466296,0.0003332985,0.00077009964,0.9590251,0.024291568,0.0026762828,0.006889029,0.00011598338],"about_ca_topic_score_codex":0.019033944,"about_ca_topic_score_gemma":0.019586848,"teacher_disagreement_score":0.019033944,"about_ca_system_score_codex":0.0018487269,"about_ca_system_score_gemma":0.0020172093,"threshold_uncertainty_score":0.048860848},"labels":[],"label_agreement":null},{"id":"W4416111891","doi":"10.48550/arxiv.2503.03862","title":"Not-Just-Scaling Laws: Towards a Better Understanding of the Downstream Impact of Language Model Design Decisions","year":2025,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; National Institute of Advanced Industrial Science and Technology; National Science Foundation","keywords":"Downstream (manufacturing); Language model; Scale (ratio); Data modeling; Code (set theory); Model building; Training (meteorology); Simulation modeling","score_opus":0.21355805319928395,"score_gpt":0.25283134413420605,"score_spread":0.0392732909349221,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416111891","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3862223,0.0045100763,0.56820834,0.011380854,0.00053855777,0.0002780429,0.0033536088,0.010277699,0.015230555],"genre_scores_gemma":[0.78972113,0.0015152871,0.19604893,0.0016664565,0.00018018477,0.0003273661,0.0037308298,0.00389616,0.0029135977],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99566925,0.0017215493,0.00021766915,0.0012120885,0.0008812323,0.00029825623],"domain_scores_gemma":[0.96683425,0.020412995,0.0014130918,0.008210169,0.0025103095,0.00061918725],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012431751,0.0020390071,0.0011084308,0.0013150839,0.00072348176,0.00534481,0.0023633768,0.0015465396,0.00277061],"category_scores_gemma":[0.06464187,0.0012133874,0.0014074134,0.0014288899,0.0023836156,0.016667463,0.0032032887,0.0072683142,0.0023613945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010753784,0.0006450482,0.10777871,0.0014301151,0.0011035753,0.0005574764,0.0029148161,0.42565846,0.04703962,0.07091201,0.03434954,0.3065352],"study_design_scores_gemma":[0.00008283844,0.0003299043,0.010668934,0.00023858485,0.00022594082,0.00015222264,0.0005898666,0.8390755,0.024162669,0.11059409,0.013744336,0.00013515443],"about_ca_topic_score_codex":0.005603741,"about_ca_topic_score_gemma":0.009801719,"teacher_disagreement_score":0.98756826,"about_ca_system_score_codex":0.0014583013,"about_ca_system_score_gemma":0.001920801,"threshold_uncertainty_score":0.06574619},"labels":[],"label_agreement":null},{"id":"W4416126051","doi":"10.1007/978-981-95-3061-8_1","title":"Global Discovery: A Global Graph-RAG Approach for Query-Focused Multimodal Summarization Across Multiple PDF Papers","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Automatic summarization; Focus (optics); Field (mathematics); Process (computing); Natural language generation; Comprehension","score_opus":0.014361331817713006,"score_gpt":0.2554453064310938,"score_spread":0.2410839746133808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416126051","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007244493,0.002004285,0.9629364,0.0005199484,0.00029749022,0.00033993085,0.0033855117,0.018726673,0.0045452304],"genre_scores_gemma":[0.11544611,0.0017443205,0.84188753,0.00048046123,0.0006124423,0.0004629564,0.015002828,0.003194207,0.021169107],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980178,0.0004752479,0.00011565023,0.00065025256,0.00058793515,0.0001531163],"domain_scores_gemma":[0.9971603,0.0010794419,0.0001999015,0.0007715928,0.0006450829,0.00014364009],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023341952,0.0021746256,0.0024700335,0.007687197,0.0016277381,0.0035393508,0.0034727857,0.0021031168,0.014499994],"category_scores_gemma":[0.006235518,0.0010090013,0.0018363484,0.008481462,0.0008385354,0.0049330415,0.0040714424,0.0019192055,0.007122892],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007286742,0.00028344107,0.0010918856,0.0007940728,0.00038911833,0.00035720857,0.00053104066,0.027799616,0.021928359,0.014173256,0.076421775,0.8555015],"study_design_scores_gemma":[0.000139042,0.0004453134,0.0022167913,0.00013192622,0.0005515936,0.0004938953,0.0008787217,0.8230491,0.021855077,0.07120478,0.07886135,0.00017230742],"about_ca_topic_score_codex":0.0058209444,"about_ca_topic_score_gemma":0.0117723,"teacher_disagreement_score":0.014499994,"about_ca_system_score_codex":0.0008539517,"about_ca_system_score_gemma":0.0017218024,"threshold_uncertainty_score":0.048507333},"labels":[],"label_agreement":null},{"id":"W4416159353","doi":"10.48550/arxiv.2511.07396","title":"C3PO: Optimized Large Language Model Cascades with Probabilistic Cost Constraints for Reasoning","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Probabilistic logic; Regret; Inference; Generalization; Scalability; Cascade; Set (abstract data type)","score_opus":0.04827555444929198,"score_gpt":0.3026207778427652,"score_spread":0.2543452233934732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416159353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022505507,0.0007048619,0.95734835,0.0008450345,0.00012909684,0.00027612396,0.0007271118,0.012642452,0.0048213955],"genre_scores_gemma":[0.4088409,0.00047934084,0.57519674,0.0011328172,0.00023225034,0.0006280838,0.003629753,0.0022587322,0.0076013114],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99830186,0.00041752853,0.00006977795,0.00050444424,0.00055165566,0.00015476545],"domain_scores_gemma":[0.9966628,0.0019809862,0.00019234382,0.00061813113,0.00038468398,0.0001610396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002665958,0.0022004617,0.0013984105,0.00086378644,0.000880626,0.0014362322,0.003973087,0.0021605303,0.0066049537],"category_scores_gemma":[0.0122356955,0.0011509598,0.0016677942,0.00083104067,0.001075534,0.003609237,0.0027021347,0.0041950885,0.0022591003],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003011951,0.00033625046,0.0014674985,0.00031580697,0.000181416,0.00021637886,0.00013123301,0.78382564,0.006034227,0.025475582,0.028209336,0.1535054],"study_design_scores_gemma":[0.000012994784,0.000016167674,0.0000488981,0.0000046013624,0.000007841743,0.000011870573,0.0000042693323,0.9916284,0.00074755325,0.006963469,0.000549283,0.0000046239684],"about_ca_topic_score_codex":0.012195173,"about_ca_topic_score_gemma":0.026741458,"teacher_disagreement_score":0.012195173,"about_ca_system_score_codex":0.0023997242,"about_ca_system_score_gemma":0.0031039226,"threshold_uncertainty_score":0.024248362},"labels":[],"label_agreement":null},{"id":"W4416181902","doi":"10.1038/s41598-025-23461-6","title":"Development of a machine learning model for automatic data extraction from breast cancer pathology reports","year":2025,"lang":"en","type":"article","venue":"Scientific Reports","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of British Columbia; University of British Columbia Hospital","funders":"Faculty of Medicine, University of British Columbia","keywords":"Pipeline (software); Scalability; Data extraction; Breast cancer; Human breast; Information extraction; Big data; Labeled data","score_opus":0.04661846175580905,"score_gpt":0.3170582380593819,"score_spread":0.27043977630357285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416181902","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04909987,0.00064721634,0.9344628,0.0015976742,0.00010325581,0.00054044777,0.00324642,0.008937111,0.0013651239],"genre_scores_gemma":[0.26974475,0.0005106373,0.7165215,0.00048040535,0.00011905356,0.0009202216,0.008536566,0.00020107876,0.0029658664],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914336,0.00022567461,0.00011110185,0.00024368567,0.00021774521,0.000058344234],"domain_scores_gemma":[0.9953945,0.003142381,0.00023423367,0.00023613135,0.00091030024,0.00008239244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039152754,0.0007461047,0.00069210556,0.0019749468,0.00052112964,0.001549393,0.0016790535,0.0011231904,0.0011764118],"category_scores_gemma":[0.0093524745,0.0004632039,0.0012426759,0.0013594837,0.00032487104,0.0014864554,0.00071213854,0.0014499977,0.0011938262],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025707672,0.0005196063,0.028857892,0.0005454335,0.00022611278,0.00043708336,0.00050180196,0.3384857,0.01668238,0.0067510377,0.019041555,0.5876944],"study_design_scores_gemma":[0.00001058299,0.00003384084,0.0011791884,0.000021505173,0.000025533274,0.00005941813,0.00003557193,0.9893985,0.0040826895,0.0022989933,0.0028405327,0.000013584902],"about_ca_topic_score_codex":0.01795025,"about_ca_topic_score_gemma":0.026631683,"teacher_disagreement_score":0.01795025,"about_ca_system_score_codex":0.0016837481,"about_ca_system_score_gemma":0.0037075644,"threshold_uncertainty_score":0.0356915},"labels":[],"label_agreement":null},{"id":"W4416183196","doi":"10.1109/gaclm67198.2025.11231924","title":"Hallucinations in Abstractive Text Summarization with Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Language model; Text graph; Source text; Natural language; Deep learning; Multi-document summarization","score_opus":0.012887607117468175,"score_gpt":0.2638264189685486,"score_spread":0.2509388118510804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416183196","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07729053,0.004388495,0.8996548,0.0010161087,0.00029033897,0.00025543247,0.0019181434,0.012958048,0.0022280782],"genre_scores_gemma":[0.5143996,0.0019242009,0.46213737,0.0006579361,0.00068781857,0.00042056778,0.011516711,0.0007435319,0.007512286],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.998844,0.0004919727,0.0001031524,0.00028211394,0.0002025252,0.000076225595],"domain_scores_gemma":[0.9968009,0.0018982022,0.0003197745,0.00042226905,0.00045082835,0.00010806908],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019213523,0.0015665665,0.00094432343,0.0015330032,0.0005016541,0.0014863167,0.0012119262,0.0010256408,0.0016803561],"category_scores_gemma":[0.0077474467,0.00042069695,0.0009074926,0.001124118,0.00057883054,0.003137063,0.0014586637,0.0016265321,0.001515564],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011352572,0.000315583,0.001955523,0.0009113978,0.00031555197,0.0006428386,0.0012624288,0.11600161,0.040098295,0.004725833,0.017475987,0.8151597],"study_design_scores_gemma":[0.00017955067,0.00049072143,0.0015056321,0.00006385801,0.00016545526,0.00020361866,0.0005065171,0.9364473,0.031149453,0.01576469,0.013450553,0.0000726581],"about_ca_topic_score_codex":0.0028436033,"about_ca_topic_score_gemma":0.0049275346,"teacher_disagreement_score":0.0028436033,"about_ca_system_score_codex":0.0005977568,"about_ca_system_score_gemma":0.0006755563,"threshold_uncertainty_score":0.010161221},"labels":[],"label_agreement":null},{"id":"W4416183388","doi":"10.1109/gaclm67198.2025.11231867","title":"Interactive Campus Crime Data Analytics: A Hybrid LLM-RAG System with Query Routing","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Routing (electronic design automation); Router; Natural language generation; Benchmark (surveying); Narrative; Natural language; Semantic data model; Work (physics)","score_opus":0.03629190059892167,"score_gpt":0.285497547253093,"score_spread":0.24920564665417133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416183388","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043285444,0.000575284,0.54971975,0.0016556784,0.00017662032,0.0010134567,0.008343599,0.3899015,0.00532863],"genre_scores_gemma":[0.34715235,0.00030641424,0.6140179,0.0016185292,0.00010956725,0.0008556619,0.023875518,0.0061687147,0.0058953413],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983595,0.00050755875,0.00015310606,0.0005718908,0.00032450072,0.000083445484],"domain_scores_gemma":[0.99684626,0.0016859649,0.00014811009,0.0008732111,0.00028407818,0.0001623731],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025053418,0.0014709991,0.0008823226,0.0016481127,0.0006563034,0.002231492,0.0029829934,0.0018114806,0.008855119],"category_scores_gemma":[0.009733217,0.00056574785,0.0009718233,0.0010295337,0.00070714176,0.0041872594,0.0048675537,0.0016319458,0.004368463],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027235725,0.001351902,0.0105506275,0.0014239671,0.00048526007,0.001801657,0.0059404043,0.04665722,0.06303241,0.020514274,0.19055296,0.6549657],"study_design_scores_gemma":[0.0003504196,0.00044556617,0.002125458,0.00007538034,0.00015280038,0.0006168997,0.0015460436,0.83366704,0.035980694,0.03182117,0.09302325,0.00019529511],"about_ca_topic_score_codex":0.0039517437,"about_ca_topic_score_gemma":0.004285029,"teacher_disagreement_score":0.008855119,"about_ca_system_score_codex":0.0007948846,"about_ca_system_score_gemma":0.000973292,"threshold_uncertainty_score":0.02962339},"labels":[],"label_agreement":null},{"id":"W4416213807","doi":"10.1109/taslpro.2025.3633059","title":"A Multifaceted Analysis of Negative Bias in Large Language Models Through the Lens of Parametric Knowledge","year":2025,"lang":"","type":"article","venue":"IEEE Transactions on Audio Speech and Language Processing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"National Research Foundation of Korea; Seoul National University","keywords":"Context (archaeology); Parametric statistics; Semantics (computer science); Negative information; Meaning (existential); Through-the-lens metering","score_opus":0.04012878237487813,"score_gpt":0.3175817237532483,"score_spread":0.2774529413783702,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416213807","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67150843,0.0014390022,0.31568992,0.0023524885,0.000084998544,0.00021586464,0.0010973149,0.0012343668,0.0063775824],"genre_scores_gemma":[0.9677229,0.00015397956,0.030471206,0.00025694253,0.000043261574,0.00012915398,0.00067044597,0.00011769588,0.00043440526],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99024194,0.006238924,0.00040731262,0.0010572663,0.001728313,0.00032623907],"domain_scores_gemma":[0.9043117,0.07934422,0.004466364,0.00741864,0.0037417652,0.00071734877],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015913477,0.0009843399,0.0006974855,0.0015495856,0.00069978123,0.0024913645,0.0010046001,0.0010974888,0.0012446844],"category_scores_gemma":[0.09236433,0.00038172668,0.0006200044,0.0011023245,0.001887934,0.0035893875,0.0025624244,0.0029503875,0.0003357316],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022444576,0.0008613569,0.2520296,0.0017009442,0.00069575355,0.0013271695,0.013700174,0.19078094,0.057132054,0.119948685,0.012120799,0.34745806],"study_design_scores_gemma":[0.000068477995,0.0004084742,0.032318704,0.00014294474,0.00015554002,0.0005181682,0.0015476278,0.80732226,0.016581642,0.13553618,0.005284459,0.00011559305],"about_ca_topic_score_codex":0.0019135352,"about_ca_topic_score_gemma":0.0026864419,"teacher_disagreement_score":0.015913477,"about_ca_system_score_codex":0.001134325,"about_ca_system_score_gemma":0.00091335573,"threshold_uncertainty_score":0.08415949},"labels":[],"label_agreement":null},{"id":"W4416217660","doi":"10.2196/76912","title":"Named Entity Recognition for Chinese Cancer Electronic Health Records—Development and Evaluation of a Domain-Specific BERT Model: Quantitative Study","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Named-entity recognition; Breast cancer; Cancer; Clinical decision support system; Electronic health record; Cancer treatment","score_opus":0.06409140617425216,"score_gpt":0.39932318930295313,"score_spread":0.33523178312870094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416217660","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9158265,0.0044817906,0.064519696,0.0009465934,0.000246357,0.0006598605,0.005358717,0.0023592962,0.0056011644],"genre_scores_gemma":[0.9506994,0.0012722714,0.03353858,0.00010623176,0.00006450063,0.0002325215,0.012270756,0.00006355112,0.0017522351],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.996527,0.001592153,0.0002526931,0.00076420396,0.00067787024,0.00018607809],"domain_scores_gemma":[0.98776954,0.008292748,0.0004986087,0.0012484854,0.0019344796,0.0002561871],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010540463,0.0012723997,0.0007768595,0.0020829232,0.00057557866,0.0012781639,0.0016096676,0.0009856852,0.001521137],"category_scores_gemma":[0.018472213,0.00032599724,0.00081969885,0.0018688295,0.00058020366,0.0042150277,0.001180088,0.0011277137,0.0005469744],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018825452,0.0019815,0.083851874,0.0018338606,0.0008191273,0.00059196114,0.00086683827,0.26983717,0.0089515345,0.004539301,0.0235798,0.6012646],"study_design_scores_gemma":[0.00004122885,0.0004229165,0.020829512,0.000063817235,0.00021133616,0.00016131715,0.0003526161,0.9685102,0.0056596524,0.00084598217,0.0028535144,0.00004788395],"about_ca_topic_score_codex":0.03612575,"about_ca_topic_score_gemma":0.025176037,"teacher_disagreement_score":0.03612575,"about_ca_system_score_codex":0.0029137903,"about_ca_system_score_gemma":0.001652037,"threshold_uncertainty_score":0.07183093},"labels":[],"label_agreement":null},{"id":"W4416219215","doi":"10.1007/978-3-032-08603-7_6","title":"Mental Health Resource Retrieval Using Semantic Similarity and Knowledge Graphs","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in business information processing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan; University of Alberta","funders":"","keywords":"Leverage (statistics); Resource (disambiguation); Ranking (information retrieval); Bridge (graph theory); Semantic similarity; Similarity (geometry); Mental health; Learning to rank; Semantics (computer science)","score_opus":0.02060005513944425,"score_gpt":0.2691981897450871,"score_spread":0.24859813460564284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416219215","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08477569,0.0060702083,0.8771393,0.0010500669,0.00028249584,0.0004736219,0.003537857,0.00542733,0.021243472],"genre_scores_gemma":[0.5260669,0.0035779485,0.45317212,0.00029146814,0.00022524953,0.0002533447,0.008723309,0.00039635962,0.0072932323],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992943,0.00018958477,0.00008640502,0.00015898637,0.00021467144,0.000056079378],"domain_scores_gemma":[0.9990637,0.00058270077,0.000056550376,0.00011603384,0.00014702042,0.00003411945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006381852,0.00051948085,0.0008232959,0.0071663535,0.0005600183,0.002186062,0.00080199953,0.0007329517,0.0034401205],"category_scores_gemma":[0.0034229676,0.00025821806,0.0011125724,0.0070443433,0.0003740644,0.0037746907,0.0013049702,0.00053896796,0.0014075548],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003136739,0.00031707552,0.003317975,0.0005268275,0.00024461746,0.00019552516,0.00029764028,0.01877356,0.00904471,0.031291932,0.021509744,0.9141666],"study_design_scores_gemma":[0.00011512939,0.00027473018,0.010418488,0.00021055719,0.0005133921,0.0010444289,0.00093052303,0.7540398,0.017538758,0.17691079,0.037867337,0.00013600897],"about_ca_topic_score_codex":0.0064729005,"about_ca_topic_score_gemma":0.007737639,"teacher_disagreement_score":0.0071663535,"about_ca_system_score_codex":0.0008418495,"about_ca_system_score_gemma":0.00084754714,"threshold_uncertainty_score":0.012870431},"labels":[],"label_agreement":null},{"id":"W4416251585","doi":"10.1109/ijcnn64981.2025.11228044","title":"Training Dynamics of a 1.7B LLaMa Model: A Data-Efficient Approach","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs","keywords":"Training (meteorology); Dynamics (music); Stability (learning theory); Sample (material); Training set; Qualitative property; Qualitative research; Variety (cybernetics)","score_opus":0.11681077259702342,"score_gpt":0.3068613338009837,"score_spread":0.19005056120396024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416251585","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16565895,0.0013991549,0.75825393,0.0050407327,0.00045514738,0.0004102213,0.0035473765,0.05591439,0.00932011],"genre_scores_gemma":[0.61074686,0.00036656373,0.36723673,0.0016905463,0.00013079497,0.0009887523,0.0061973636,0.0057178307,0.006924662],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985067,0.00057605974,0.000082428065,0.0004704446,0.00022951954,0.00013485437],"domain_scores_gemma":[0.99373627,0.003970453,0.00013513352,0.0011695118,0.00077593257,0.00021263279],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032434636,0.0015380995,0.0012162363,0.0007161634,0.0009533482,0.0026992026,0.0035855763,0.00190443,0.007330558],"category_scores_gemma":[0.022256877,0.0014604861,0.0012813952,0.0006250732,0.0012639873,0.0053789294,0.0029454373,0.0055054817,0.004749338],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012845302,0.00045152742,0.007612625,0.0004614625,0.0002471091,0.00036434646,0.0015894372,0.7105737,0.014510924,0.013615715,0.032123715,0.21716492],"study_design_scores_gemma":[0.000038117203,0.000060157345,0.00022910227,0.00002804027,0.000019731786,0.000031347477,0.000117550415,0.98512,0.0025943103,0.0083378665,0.0034072504,0.000016411695],"about_ca_topic_score_codex":0.013171269,"about_ca_topic_score_gemma":0.02562619,"teacher_disagreement_score":0.013171269,"about_ca_system_score_codex":0.0019473252,"about_ca_system_score_gemma":0.0028465157,"threshold_uncertainty_score":0.026189208},"labels":[],"label_agreement":null},{"id":"W4416266029","doi":"10.1016/j.neunet.2025.108218","title":"Optimizing boundary dynamics for nested named entity recognition via semantic refinement and trimming","year":2025,"lang":"en","type":"article","venue":"Neural Networks","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"National Key Research and Development Program of China; Key Technologies Research and Development Program; Natural Science Foundation of Sichuan Province; National Natural Science Foundation of China","keywords":"Ambiguity; Trimming; Entity linking; Kernel (algebra); Construct (python library); Semantics (computer science); Boundary (topology); Semantic data model; Task (project management)","score_opus":0.017689362879480426,"score_gpt":0.2488046597694865,"score_spread":0.2311152968900061,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416266029","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02031361,0.00023462594,0.97591156,0.00019107107,0.00005344066,0.000041527364,0.00012667807,0.0021580213,0.00096938625],"genre_scores_gemma":[0.51550037,0.0002612544,0.47611877,0.00032577294,0.00009434588,0.000271594,0.0013868046,0.0011009632,0.0049401238],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915934,0.00016435138,0.0000540576,0.0003748203,0.0001342878,0.000113200564],"domain_scores_gemma":[0.99754864,0.0015882152,0.00013962299,0.0003055466,0.0002970369,0.000120941186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015928623,0.0011240606,0.0020590052,0.0012279314,0.0010713602,0.0015618013,0.0028567202,0.0027737883,0.0048085316],"category_scores_gemma":[0.006969248,0.0012776344,0.001656593,0.0013630651,0.0011538408,0.004592983,0.0028015664,0.0031406647,0.0018319468],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038791817,0.00016279506,0.0010508639,0.00016838132,0.00009016216,0.0001895959,0.0003165856,0.70173854,0.012271874,0.017885722,0.0056614107,0.26007608],"study_design_scores_gemma":[0.000004069441,0.0000062961026,0.00003683352,0.000003529325,0.0000041058383,0.0000064507785,0.000013108009,0.99460113,0.0006362797,0.00450527,0.00017972046,0.0000031747074],"about_ca_topic_score_codex":0.014336491,"about_ca_topic_score_gemma":0.022187905,"teacher_disagreement_score":0.014336491,"about_ca_system_score_codex":0.0012434793,"about_ca_system_score_gemma":0.0018322783,"threshold_uncertainty_score":0.0285061},"labels":[],"label_agreement":null},{"id":"W4416372001","doi":"10.1162/tacl.a.50","title":"Objectifying the Subjective: Cognitive Biases in Topic Interpretations","year":2025,"lang":"en","type":"article","venue":"Transactions of the Association for Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Representativeness heuristic; Interpretation (philosophy); Salient; Cognitive bias; Quality (philosophy); Heuristics; Context (archaeology); Coherence (philosophical gambling strategy); Cognition","score_opus":0.022742452538391057,"score_gpt":0.3057547969801711,"score_spread":0.28301234444178003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416372001","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43459925,0.002171145,0.5183964,0.007792767,0.0002877057,0.0005416407,0.00019679921,0.0008541177,0.035160225],"genre_scores_gemma":[0.9612249,0.00026422227,0.036805496,0.000507564,0.0000874319,0.00020316124,0.000067741756,0.0001784945,0.00066097564],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.88410014,0.09107905,0.0042135315,0.005254555,0.014171165,0.001181616],"domain_scores_gemma":[0.60282075,0.32326096,0.025631877,0.026621528,0.018907746,0.0027571463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08430993,0.00093414023,0.00075445813,0.0036769959,0.0023747298,0.01286434,0.0015870453,0.0018100016,0.0025265408],"category_scores_gemma":[0.33018917,0.0008402878,0.00092771766,0.002216997,0.012971289,0.0135213705,0.0058147092,0.0033692953,0.00045303325],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009672212,0.00023474132,0.07207017,0.0017277579,0.00048943283,0.00051189255,0.3134026,0.007408341,0.014791591,0.33517653,0.0067617763,0.246458],"study_design_scores_gemma":[0.00017903977,0.00047290372,0.04222086,0.0015670885,0.0003717869,0.0008689147,0.07332376,0.060347058,0.013971278,0.76445043,0.04163784,0.0005890975],"about_ca_topic_score_codex":0.0021105472,"about_ca_topic_score_gemma":0.0014908286,"teacher_disagreement_score":0.08430993,"about_ca_system_score_codex":0.0027949854,"about_ca_system_score_gemma":0.0018443923,"threshold_uncertainty_score":0.44587886},"labels":[],"label_agreement":null},{"id":"W4416408255","doi":"10.48550/arxiv.2507.19219","title":"How Much Do Large Language Model Cheat on Evaluation? Benchmarking Overestimation under the One-Time-Pad-Based Framework","year":2025,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Innovation and Technology Fund; Institute for Catastrophic Loss Reduction; University of Science and Technology of China; Microsoft Research; National Natural Science Foundation of China; Hong Kong Polytechnic University","keywords":"Benchmarking; Benchmark (surveying); Key (lock); Quality (philosophy); Test (biology); Measure (data warehouse)","score_opus":0.08765754670620739,"score_gpt":0.2418694500749871,"score_spread":0.1542119033687797,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416408255","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41954243,0.010756451,0.45698652,0.014636916,0.0025000346,0.0009993317,0.004449752,0.053536456,0.03659216],"genre_scores_gemma":[0.8490027,0.00070732436,0.13598238,0.0023598862,0.00028641542,0.00051194034,0.004140669,0.0039802864,0.0030284217],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.92967695,0.043169655,0.0044753896,0.0052123936,0.014977687,0.0024878932],"domain_scores_gemma":[0.871769,0.06280755,0.0050995955,0.046224426,0.011279456,0.0028198494],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.040346805,0.0018931546,0.0016587632,0.0020514464,0.0010722807,0.0049859206,0.0041479254,0.002337679,0.0039903047],"category_scores_gemma":[0.19115554,0.0008003547,0.0012461084,0.0020105192,0.0033741645,0.009788502,0.0057066567,0.004138034,0.0032142468],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039669457,0.0012617862,0.051041797,0.0022954685,0.0010496159,0.0007150833,0.0016090174,0.17845243,0.029894367,0.07987857,0.10961937,0.54021543],"study_design_scores_gemma":[0.0007776256,0.0026004694,0.010031555,0.00065992516,0.00024431464,0.00095153594,0.0008783797,0.769046,0.056513343,0.118779205,0.039259117,0.00025850066],"about_ca_topic_score_codex":0.0026412746,"about_ca_topic_score_gemma":0.0038083147,"teacher_disagreement_score":0.9596532,"about_ca_system_score_codex":0.0029869883,"about_ca_system_score_gemma":0.0047531263,"threshold_uncertainty_score":0.21337682},"labels":[],"label_agreement":null},{"id":"W4416414115","doi":"10.1016/j.neucom.2025.132127","title":"Document-level event argument extraction by leveraging span boundaries and argument dependencies","year":2025,"lang":"en","type":"article","venue":"Neurocomputing","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Yunnan Provincial Science and Technology Department; Key Laboratory of Software Engineering of Yunnan Province; Yunnan Provincial Department of Education","keywords":"Discriminative model; Argument (complex analysis); Span (engineering); Event (particle physics); Context (archaeology); Matching (statistics); Representation (politics); Dependency graph","score_opus":0.019385358476909205,"score_gpt":0.27156954284096224,"score_spread":0.25218418436405304,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416414115","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11198003,0.008907751,0.8030356,0.0027314741,0.0012815086,0.0010411341,0.02858851,0.017055022,0.02537905],"genre_scores_gemma":[0.52451575,0.00337913,0.40071046,0.00038491416,0.0016085592,0.00063215924,0.056336634,0.0015570933,0.010875293],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982522,0.0002945293,0.0002673525,0.00056942133,0.000442132,0.00017430159],"domain_scores_gemma":[0.99361044,0.0035544904,0.00052494253,0.00048483082,0.001557473,0.00026777887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016383429,0.0018037527,0.0011026616,0.010876134,0.0013374857,0.003282412,0.0012532246,0.0022964333,0.008955954],"category_scores_gemma":[0.009788635,0.00061991403,0.001746304,0.006438319,0.00038877578,0.005542758,0.0022680126,0.0024235777,0.008495491],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010555115,0.0006053482,0.019122424,0.0025423893,0.00035423008,0.001687626,0.0015114429,0.006739829,0.050830554,0.015809687,0.0748983,0.8248427],"study_design_scores_gemma":[0.00027726917,0.0004837291,0.04153983,0.0011444737,0.001731838,0.0032901228,0.0026563776,0.58390987,0.079767235,0.08882086,0.19609597,0.00028244068],"about_ca_topic_score_codex":0.0023208295,"about_ca_topic_score_gemma":0.0044385605,"teacher_disagreement_score":0.010876134,"about_ca_system_score_codex":0.00074083416,"about_ca_system_score_gemma":0.0022004272,"threshold_uncertainty_score":0.029960632},"labels":[],"label_agreement":null},{"id":"W4416445260","doi":"10.48550/arxiv.2505.03563","title":"Say It Another Way: Auditing LLMs with a User-Grounded Automated Paraphrasing Framework","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fonds de recherche du Québec – Nature et technologies; Compute Canada; Canadian Institute for Advanced Research","keywords":"Audit; Quality (philosophy); Natural language; Strengths and weaknesses; Natural (archaeology); Language model; Semantics (computer science)","score_opus":0.045067601921034116,"score_gpt":0.29144446469529406,"score_spread":0.24637686277425994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416445260","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040687572,0.0007786467,0.9202836,0.0021533803,0.00013007542,0.00051942165,0.0028644255,0.029316165,0.0032666188],"genre_scores_gemma":[0.4327767,0.0002694856,0.55762815,0.0008951072,0.00007344828,0.0004721661,0.004361213,0.0018103152,0.0017133976],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9759754,0.016764829,0.0012647896,0.0023848743,0.0031575235,0.00045266654],"domain_scores_gemma":[0.9462317,0.023399938,0.00299733,0.021948325,0.0047459826,0.0006767896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013392629,0.00118855,0.0008916912,0.0021411579,0.00086977053,0.004978543,0.0026709118,0.0021854423,0.003099679],"category_scores_gemma":[0.087731056,0.00069686864,0.0010057623,0.0017009936,0.0019726935,0.007490433,0.005791174,0.0036712391,0.0028841386],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014052641,0.00081539043,0.02164625,0.0019354668,0.00039021883,0.0008890699,0.009292052,0.08418111,0.049798742,0.05850671,0.049694158,0.7214456],"study_design_scores_gemma":[0.00017828959,0.0005853031,0.0040794318,0.00033969374,0.00010696595,0.0006421484,0.0019360601,0.78157216,0.04441573,0.118679404,0.0472629,0.00020194767],"about_ca_topic_score_codex":0.0040875785,"about_ca_topic_score_gemma":0.006987161,"teacher_disagreement_score":0.013392629,"about_ca_system_score_codex":0.0011091697,"about_ca_system_score_gemma":0.0032286458,"threshold_uncertainty_score":0.07082784},"labels":[],"label_agreement":null},{"id":"W4416575445","doi":"10.3390/computers14120508","title":"NewsSumm: The World’s Largest Human-Annotated Multi-Document News Summarization Dataset for Indian English","year":2025,"lang":"en","type":"article","venue":"Computers","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Automatic summarization; Newspaper; Timeline; Journalism; Consistency (knowledge bases); Scale (ratio); Indian English; Domain (mathematical analysis); Quality (philosophy)","score_opus":0.02732516811302219,"score_gpt":0.3085124555013454,"score_spread":0.28118728738832316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416575445","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015126058,0.0025592372,0.004601604,0.0008140111,0.00051511766,0.000333671,0.94692576,0.013187622,0.01593684],"genre_scores_gemma":[0.00906199,0.00040766533,0.008968535,0.00014451315,0.00010808332,0.00030192922,0.9769732,0.0004865294,0.0035475674],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99826425,0.0003768541,0.00030089333,0.00040863376,0.0004999057,0.00014949689],"domain_scores_gemma":[0.99624985,0.00094311556,0.00045651733,0.0006662777,0.0013643266,0.00031984563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012909626,0.0015208198,0.00075530517,0.0100604575,0.0016562681,0.0019296566,0.0016939581,0.0011259693,0.015234121],"category_scores_gemma":[0.007010464,0.00034246742,0.00081451953,0.008165689,0.000635856,0.0017441906,0.0017561602,0.0012680315,0.017066015],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035985874,0.00012329957,0.0042235833,0.004117615,0.00013581477,0.00046839254,0.0014094557,0.0008871929,0.0076354193,0.0015155062,0.90881366,0.070310175],"study_design_scores_gemma":[0.00012918025,0.000111515605,0.029372143,0.00047950807,0.00018191745,0.00055220677,0.0017202413,0.0044653807,0.0078286985,0.0012129287,0.9537964,0.00014997309],"about_ca_topic_score_codex":0.02865908,"about_ca_topic_score_gemma":0.06357514,"teacher_disagreement_score":0.02865908,"about_ca_system_score_codex":0.0012539242,"about_ca_system_score_gemma":0.0022587616,"threshold_uncertainty_score":0.056984544},"labels":[],"label_agreement":null},{"id":"W4416617284","doi":"10.1109/icspis68676.2025.11551719","title":"Enhancing Reasoning Skills in Small Persian Medical Language Models Can Outperform Large-Scale Data Training","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Persian; Baseline (sea); Language model; Training set; Preference; Verbal reasoning; Training (meteorology); Variety (cybernetics)","score_opus":0.03256544562111078,"score_gpt":0.2876779793084234,"score_spread":0.2551125336873126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416617284","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46678463,0.0035004586,0.46199483,0.003937759,0.0006485497,0.0005406323,0.0028013818,0.039682474,0.020109285],"genre_scores_gemma":[0.7995741,0.0003813988,0.18736675,0.0014521412,0.00007866839,0.00019896166,0.005088282,0.00052131707,0.0053385133],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9988943,0.00043453087,0.00007658881,0.00034282883,0.00016105988,0.00009069552],"domain_scores_gemma":[0.9961456,0.0026574621,0.00014343865,0.00051957247,0.0003868384,0.00014699741],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002572019,0.0014348378,0.00055363943,0.00044788758,0.00035442231,0.000899947,0.0014991352,0.0010425871,0.005554588],"category_scores_gemma":[0.010133161,0.00040780436,0.0008954548,0.00035200795,0.0006736138,0.0017707859,0.0014601578,0.0027698223,0.002538916],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012771445,0.00081201806,0.012352458,0.0008035625,0.0002283828,0.0004006439,0.00050511124,0.3568689,0.026112795,0.0039808,0.022565939,0.57409227],"study_design_scores_gemma":[0.00014242307,0.0003433013,0.0019228938,0.0000648526,0.00006266241,0.0001172725,0.00017412253,0.96908075,0.015477831,0.005682075,0.006902523,0.000029242845],"about_ca_topic_score_codex":0.004959032,"about_ca_topic_score_gemma":0.0120189795,"teacher_disagreement_score":0.005554588,"about_ca_system_score_codex":0.00081954774,"about_ca_system_score_gemma":0.0019134635,"threshold_uncertainty_score":0.018581927},"labels":[],"label_agreement":null},{"id":"W4416673024","doi":"10.1007/s10462-025-11421-5","title":"Exploring unanswerability in machine reading comprehension: approaches, benchmarks, and open challenges","year":2025,"lang":"en","type":"article","venue":"Artificial Intelligence Review","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; Ted Rogers Centre for Heart Research; University of Guelph","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Key (lock); Comprehension; Reading (process); Work (physics); Reading comprehension","score_opus":0.5294704198234306,"score_gpt":0.3733258316109949,"score_spread":0.15614458821243565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416673024","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12405238,0.36126664,0.43430424,0.028307877,0.0010516283,0.0009868178,0.0075694425,0.007834865,0.03462609],"genre_scores_gemma":[0.5911626,0.08470377,0.29520798,0.0027643067,0.0017862888,0.0015923325,0.017641155,0.001239011,0.0039025715],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97874326,0.012601821,0.0015536756,0.0031245374,0.0035594814,0.00041716735],"domain_scores_gemma":[0.8143226,0.15807508,0.005269864,0.0081382785,0.012861979,0.0013321636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023992743,0.0023046548,0.0023744216,0.01101005,0.0011842686,0.008471979,0.003612074,0.0025123488,0.004197866],"category_scores_gemma":[0.11846011,0.00093904376,0.0017880958,0.009437413,0.0026787515,0.018632,0.0044889194,0.0045526032,0.002615258],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026956655,0.00033671406,0.030464925,0.0112334555,0.00077127886,0.000167359,0.009716607,0.010498149,0.003356868,0.02014232,0.024622578,0.88842016],"study_design_scores_gemma":[0.00016014285,0.00090327807,0.10677756,0.012450266,0.0013422517,0.001412508,0.027744748,0.20567831,0.01455938,0.37800696,0.25023174,0.00073281105],"about_ca_topic_score_codex":0.007243303,"about_ca_topic_score_gemma":0.008748617,"teacher_disagreement_score":0.023992743,"about_ca_system_score_codex":0.002458365,"about_ca_system_score_gemma":0.0037535501,"threshold_uncertainty_score":0.12688726},"labels":[],"label_agreement":null},{"id":"W4416677522","doi":"10.1109/models67397.2025.00024","title":"SHERPA: A Model-Driven Framework for Large Language Model Execution","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Variety (cybernetics); State (computer science); Class (philosophy); Structuring; Best practice; Baseline (sea); Language model","score_opus":0.02454095306884229,"score_gpt":0.3090076620587478,"score_spread":0.2844667089899055,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416677522","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011262128,0.00014409139,0.9692402,0.00016103679,0.00004977396,0.00014575904,0.00051243807,0.028012063,0.00060849375],"genre_scores_gemma":[0.06283162,0.0003399174,0.92641664,0.00030524444,0.00006748416,0.0008991231,0.0032246963,0.0043920884,0.0015232727],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975968,0.0010497607,0.00018572317,0.00048294317,0.00055939745,0.00012533326],"domain_scores_gemma":[0.9936354,0.0044604386,0.00029038097,0.0008225817,0.00057484716,0.00021636547],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004209899,0.002220729,0.0012304924,0.0014806011,0.0009234806,0.0028184848,0.0048357877,0.0018572437,0.008848342],"category_scores_gemma":[0.015046536,0.0018689957,0.00327005,0.0009887528,0.0012720968,0.0031667796,0.0034141869,0.0050269184,0.004436016],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004730589,0.00028556786,0.002399218,0.00085605716,0.00041301778,0.00040882739,0.00060573657,0.6872778,0.007122461,0.07720851,0.033740003,0.18920983],"study_design_scores_gemma":[0.000028598875,0.000017307015,0.000037412967,0.0000149534535,0.000013697885,0.000018327384,0.000012281328,0.972329,0.0011479249,0.02121806,0.005148024,0.0000143736725],"about_ca_topic_score_codex":0.011124986,"about_ca_topic_score_gemma":0.0247368,"teacher_disagreement_score":0.011124986,"about_ca_system_score_codex":0.0015351715,"about_ca_system_score_gemma":0.0048248037,"threshold_uncertainty_score":0.02960062},"labels":[],"label_agreement":null},{"id":"W4416677557","doi":"10.1109/models67397.2025.00018","title":"Accurate and Consistent Graph Model Generation from Text with Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada); McGill University","funders":"","keywords":"Graph; Language model; Probabilistic logic; Constraint (computer-aided design); Consistency (knowledge bases); Graphical model; Syntax","score_opus":0.03510477081533805,"score_gpt":0.258425097175378,"score_spread":0.22332032636003993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416677557","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02775217,0.00068189576,0.92882997,0.001026596,0.00014234975,0.00038534534,0.0065723476,0.03280847,0.0018008917],"genre_scores_gemma":[0.17794716,0.0005889724,0.76846355,0.00052187627,0.00007007824,0.0005152321,0.045980014,0.0036384452,0.002274719],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99670213,0.0012964624,0.00021494256,0.0008031438,0.0008714838,0.000111774134],"domain_scores_gemma":[0.98517656,0.0099707935,0.000588428,0.0025635592,0.0015328252,0.00016787655],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032495277,0.0019534978,0.0009713674,0.003240014,0.0008513603,0.0019525031,0.0030270633,0.0019117547,0.0028227405],"category_scores_gemma":[0.024457026,0.00069474895,0.0028057697,0.00262293,0.0008165735,0.004006482,0.0021968463,0.002170433,0.002058194],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036836724,0.00038013168,0.0071953973,0.00195993,0.00047219332,0.0014681409,0.0010940799,0.45740768,0.014713001,0.027344331,0.06345014,0.42414656],"study_design_scores_gemma":[0.00007386018,0.00003767268,0.00043264806,0.00005997385,0.00007124584,0.00017471678,0.00016833175,0.945741,0.0068798726,0.034385026,0.011940364,0.00003526104],"about_ca_topic_score_codex":0.0110875,"about_ca_topic_score_gemma":0.026631169,"teacher_disagreement_score":0.0110875,"about_ca_system_score_codex":0.0015470057,"about_ca_system_score_gemma":0.0027190023,"threshold_uncertainty_score":0.02204597},"labels":[],"label_agreement":null},{"id":"W4416679312","doi":"10.48550/arxiv.2504.05058","title":"Not All Data Are Unlearned Equally","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Samsung; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Task (project management); Context (archaeology); Phone; Training set; Data collection","score_opus":0.25539735961457605,"score_gpt":0.34824701290357185,"score_spread":0.0928496532889958,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416679312","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29611292,0.0039959773,0.67425126,0.006190428,0.0005820304,0.0004726908,0.0027580976,0.004920246,0.010716327],"genre_scores_gemma":[0.8769142,0.0007958804,0.11287613,0.0017268391,0.00020198655,0.00024315529,0.0044798963,0.0006274614,0.0021344433],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97762954,0.011049288,0.0016757054,0.0045365393,0.004390032,0.00071898673],"domain_scores_gemma":[0.8845128,0.0722556,0.004393446,0.032328658,0.00503204,0.0014775157],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019039597,0.001775406,0.0017118362,0.0015728491,0.0013890332,0.004474514,0.0024680297,0.002966379,0.002402724],"category_scores_gemma":[0.119005874,0.00092381286,0.0017007686,0.0017489842,0.0032294686,0.01490502,0.005251132,0.006922708,0.0014028243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023514684,0.0008445127,0.0688672,0.0014742293,0.0009938471,0.0006401901,0.0021961448,0.21481693,0.011137697,0.04038132,0.015913643,0.6403828],"study_design_scores_gemma":[0.00018971984,0.00076175644,0.012513231,0.0004971188,0.00039497667,0.0010782292,0.0011843126,0.69859767,0.030570416,0.23243393,0.021586671,0.00019197066],"about_ca_topic_score_codex":0.0030395729,"about_ca_topic_score_gemma":0.0039751725,"teacher_disagreement_score":0.019039597,"about_ca_system_score_codex":0.0018353253,"about_ca_system_score_gemma":0.0016432969,"threshold_uncertainty_score":0.10069221},"labels":[],"label_agreement":null},{"id":"W4416747828","doi":"10.1109/intcec65580.2025.11255828","title":"Position Bias Across LLM Model Families","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Suite; Process (computing); Position (finance); Value (mathematics); Position paper","score_opus":0.05111061156713082,"score_gpt":0.3110795303689483,"score_spread":0.2599689188018175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416747828","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.65101385,0.0050901724,0.28924558,0.009489572,0.0003046842,0.0005218658,0.0051808245,0.0027554173,0.036397975],"genre_scores_gemma":[0.96098137,0.0006162994,0.028862854,0.0009931203,0.00011585556,0.00031918325,0.0038706458,0.0006166468,0.003624093],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9793617,0.014216724,0.00085162075,0.002842191,0.0020228145,0.00070506433],"domain_scores_gemma":[0.8042761,0.17460418,0.004406697,0.011240485,0.004227346,0.001245146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.036482986,0.0012579742,0.0015305418,0.0024071957,0.0018736405,0.004735243,0.002689974,0.002452013,0.011275384],"category_scores_gemma":[0.15881479,0.00081468484,0.0020125445,0.0022158446,0.0020108665,0.0072259395,0.0028815044,0.004314999,0.0022213527],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0046549793,0.0006355061,0.16476949,0.0012434012,0.001650048,0.00092263706,0.006431664,0.2931918,0.0032298728,0.29716703,0.039065853,0.18703783],"study_design_scores_gemma":[0.00037003832,0.00027073149,0.012824536,0.0002947219,0.00032946572,0.00040100908,0.0014824072,0.6873746,0.002196365,0.28047252,0.013856674,0.00012692646],"about_ca_topic_score_codex":0.012467564,"about_ca_topic_score_gemma":0.015171947,"teacher_disagreement_score":0.036482986,"about_ca_system_score_codex":0.0038720255,"about_ca_system_score_gemma":0.0017826385,"threshold_uncertainty_score":0.1929428},"labels":[],"label_agreement":null},{"id":"W4416780702","doi":"10.1016/j.esmorw.2025.100615","title":"419P Development of the user-friendly decision aid rule-based evaluation and support tool (REST) for optimizing the resources of an information extraction task","year":2025,"lang":"en","type":"article","venue":"ESMO Real World Data and Digital Oncology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Consejo Superior de Investigaciones Científicas; University of Texas MD Anderson Cancer Center; Centre hospitalier universitaire Sainte-Justine; Fundação Champalimaud; Universidad de Buenos Aires","keywords":"Task (project management); Information extraction; Development (topology); Task analysis; Natural language; Decision support system","score_opus":0.03554035092718389,"score_gpt":0.3524967936112997,"score_spread":0.3169564426841158,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416780702","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011151921,0.00010445567,0.8132939,0.0004848242,0.00012222382,0.001096652,0.0053650495,0.16253652,0.0058445283],"genre_scores_gemma":[0.08237351,0.000109949295,0.88862795,0.00041212607,0.00006652861,0.0010988482,0.0108829355,0.007521704,0.008906522],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99588335,0.0013170548,0.0005724578,0.00072139886,0.0012794462,0.00022638575],"domain_scores_gemma":[0.9837705,0.0070854807,0.0006696372,0.002509752,0.0054561705,0.00050848053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00837121,0.0019288862,0.0011378559,0.0029497351,0.0007342219,0.0035658774,0.0028891717,0.0013778481,0.02209073],"category_scores_gemma":[0.024786701,0.0008668656,0.0014885992,0.0010557878,0.000546378,0.0029656952,0.0022745533,0.0016452773,0.013267426],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015120297,0.00087044004,0.008275025,0.0014348802,0.00035454423,0.0017105531,0.0010459414,0.025544062,0.03955597,0.010506146,0.12244236,0.78674805],"study_design_scores_gemma":[0.0003377662,0.00037942082,0.004919006,0.00043159298,0.00024520123,0.0010166299,0.00033195756,0.66459715,0.16191784,0.010843053,0.15469278,0.0002876524],"about_ca_topic_score_codex":0.004828475,"about_ca_topic_score_gemma":0.004166749,"teacher_disagreement_score":0.02209073,"about_ca_system_score_codex":0.0008815091,"about_ca_system_score_gemma":0.002596922,"threshold_uncertainty_score":0.07390088},"labels":[],"label_agreement":null},{"id":"W4416924418","doi":"10.1109/indiscon66021.2025.11254534","title":"Why T5 Forgets Entities: Diagnosing and Mitigating Attention Failures in Summarization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Automatic summarization; Multi-document summarization; Salient; Key (lock)","score_opus":0.011790925603683687,"score_gpt":0.24817892907398917,"score_spread":0.23638800347030547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416924418","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48861557,0.006622338,0.42456067,0.0038061182,0.00090082304,0.0007832317,0.006041144,0.06219219,0.006477856],"genre_scores_gemma":[0.78016484,0.00095837284,0.19752972,0.0008694245,0.0004412795,0.00026259114,0.013412984,0.0013640398,0.0049967836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99872094,0.00032091458,0.00015290099,0.00038239558,0.0002428972,0.00017994487],"domain_scores_gemma":[0.99206984,0.0036287932,0.0007154502,0.00123465,0.0018777269,0.00047351344],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029399237,0.002031948,0.0009243724,0.0026180646,0.0010338993,0.0022131368,0.00192946,0.0016430793,0.0028601848],"category_scores_gemma":[0.017832179,0.0004346883,0.00073486235,0.0013141631,0.0007003473,0.0043888493,0.0023287758,0.0017484882,0.0029385653],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021997767,0.00030300563,0.041046616,0.0014256416,0.00047884308,0.0012983016,0.0026828775,0.023247883,0.10065806,0.0023040606,0.05590971,0.76844525],"study_design_scores_gemma":[0.00020898061,0.002054114,0.026431864,0.00027376565,0.0007964805,0.0012137305,0.0027143413,0.7269374,0.19175012,0.013721854,0.03372741,0.00016990522],"about_ca_topic_score_codex":0.0064212508,"about_ca_topic_score_gemma":0.010447359,"teacher_disagreement_score":0.0064212508,"about_ca_system_score_codex":0.0007420717,"about_ca_system_score_gemma":0.0011842569,"threshold_uncertainty_score":0.015547991},"labels":[],"label_agreement":null},{"id":"W4416940269","doi":"10.1016/j.imu.2025.101720","title":"SuSastho.AI: A multimodal medical copilot for adolescents using evidence-based medicine and large language models","year":2025,"lang":"en","type":"article","venue":"Informatics in Medicine Unlocked","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"International Development Research Centre; United International University; Bill and Melinda Gates Foundation","keywords":"Mental health; Correctness; Health care; Unintended consequences; Mental healthcare; Reproductive health; Face (sociological concept); Occupational safety and health; Health equity","score_opus":0.06801586580321954,"score_gpt":0.3709301582021793,"score_spread":0.30291429239895973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416940269","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4806992,0.0048082797,0.23665717,0.009223062,0.0010782034,0.009598454,0.16886182,0.054284953,0.034788862],"genre_scores_gemma":[0.477921,0.0017462217,0.37899727,0.0018832855,0.0001762953,0.0048703807,0.11686848,0.0009382817,0.016598782],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990522,0.0005511686,0.000110573164,0.00012683806,0.0001143921,0.000044736054],"domain_scores_gemma":[0.99699956,0.002176941,0.00016246142,0.00021508483,0.00025166452,0.00019435873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016370745,0.0010910861,0.0004380314,0.0013642878,0.0005023755,0.001112673,0.0010074236,0.0010510104,0.014820095],"category_scores_gemma":[0.0074258703,0.00021225982,0.0009776629,0.00061334064,0.0003027586,0.0014888069,0.0023251106,0.00075271813,0.004384936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027861868,0.0016354314,0.0316822,0.009518422,0.00038799067,0.004013424,0.0067218933,0.016921489,0.038126823,0.004823042,0.17439915,0.70898396],"study_design_scores_gemma":[0.00152052,0.0049877856,0.051679764,0.0027618716,0.0007210312,0.0046772845,0.01743208,0.33819,0.047016136,0.024195889,0.5060543,0.00076333276],"about_ca_topic_score_codex":0.002788576,"about_ca_topic_score_gemma":0.007591303,"teacher_disagreement_score":0.014820095,"about_ca_system_score_codex":0.00058477616,"about_ca_system_score_gemma":0.0011199537,"threshold_uncertainty_score":0.04957819},"labels":[],"label_agreement":null},{"id":"W4416996294","doi":"10.1007/s43681-025-00838-x","title":"Toward ethical AI through Bayesian uncertainty in neural question answering","year":2025,"lang":"en","type":"article","venue":"AI and Ethics","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HealthForceOntario","funders":"","keywords":"Interpretability; Bayesian probability; Bayesian inference; Inference; Artificial neural network; A priori and a posteriori; Calibration; Question answering","score_opus":0.05108168716002209,"score_gpt":0.35554357533227976,"score_spread":0.30446188817225767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416996294","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027731972,0.0009871862,0.9575053,0.007444353,0.00007187794,0.00006328082,0.00025303423,0.00037570624,0.0055673085],"genre_scores_gemma":[0.79003024,0.0007866109,0.2035187,0.0012373804,0.0003242999,0.0001782171,0.0006828796,0.00011117509,0.0031305011],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965024,0.0023546673,0.00013421148,0.00045672432,0.00039106762,0.0001608717],"domain_scores_gemma":[0.9688167,0.028122805,0.0007386344,0.0008956577,0.0010440646,0.0003820609],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00659778,0.00056477403,0.0010489188,0.0016082481,0.0008832276,0.0033240009,0.0018626729,0.0026855897,0.0034410036],"category_scores_gemma":[0.043395136,0.00073679356,0.0009987049,0.0012815392,0.0022831864,0.007866592,0.0032550322,0.0053159706,0.00056288025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030137677,0.00027360028,0.0047621564,0.00042909526,0.0002395387,0.00012987736,0.0018864111,0.2027911,0.0018808809,0.5623908,0.01149306,0.21342205],"study_design_scores_gemma":[0.000011229729,0.000009800126,0.00025555823,0.000028513643,0.000019189838,0.00001790659,0.000076941265,0.5276681,0.00039090097,0.4702534,0.0012574369,0.000010965744],"about_ca_topic_score_codex":0.005752949,"about_ca_topic_score_gemma":0.0058954298,"teacher_disagreement_score":0.00659778,"about_ca_system_score_codex":0.0017209272,"about_ca_system_score_gemma":0.0014003122,"threshold_uncertainty_score":0.034892857},"labels":[],"label_agreement":null},{"id":"W4417003228","doi":"10.48550/arxiv.2512.02240","title":"Lightweight Latent Reasoning for Narrative Tasks","year":2025,"lang":"","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Security token; Latent variable; Encoding (memory); Reinforcement learning; Matching (statistics); Debiasing; Language model; Reasoning system","score_opus":0.08023349451410348,"score_gpt":0.20706437581573192,"score_spread":0.12683088130162845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417003228","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012572131,0.0005173605,0.9704315,0.00040994986,0.000052882235,0.00012364217,0.00071815826,0.013442911,0.0017314226],"genre_scores_gemma":[0.36179596,0.00051137037,0.623583,0.00037971025,0.000115898896,0.00037307263,0.0051324773,0.0012966574,0.0068118856],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99860376,0.0005897588,0.00008375997,0.00037900655,0.00023485241,0.0001088038],"domain_scores_gemma":[0.9965926,0.0021165241,0.00022413624,0.0007325109,0.0001779864,0.00015619917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017894343,0.0014338347,0.0008339361,0.0009979123,0.000666879,0.001731444,0.0031887295,0.0014574467,0.009201046],"category_scores_gemma":[0.008637699,0.000792008,0.0018916449,0.00084009283,0.0007822821,0.0049200635,0.0023695582,0.0031876233,0.0044858474],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077462965,0.00052925805,0.003788977,0.0010386748,0.0002150987,0.00041324183,0.001017192,0.22737087,0.01860119,0.050335094,0.030499365,0.66541636],"study_design_scores_gemma":[0.00005785174,0.00005834747,0.0002979329,0.000032223736,0.000027781882,0.00007693755,0.00008515512,0.93552274,0.00483036,0.053462025,0.005526781,0.00002184617],"about_ca_topic_score_codex":0.0037402802,"about_ca_topic_score_gemma":0.009176453,"teacher_disagreement_score":0.009201046,"about_ca_system_score_codex":0.0012731047,"about_ca_system_score_gemma":0.0014276223,"threshold_uncertainty_score":0.030780554},"labels":[],"label_agreement":null},{"id":"W4417026595","doi":"10.1007/s10664-025-10740-z","title":"Empirical studies of parameter efficient methods for large language models of code and knowledge transfer to R","year":2025,"lang":"en","type":"article","venue":"Empirical Software Engineering","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Okanagan University College; University of British Columbia, Okanagan Campus; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Empirical research; Code (set theory); Knowledge transfer; Language model; Resource (disambiguation)","score_opus":0.06950578563380058,"score_gpt":0.4078640021553274,"score_spread":0.33835821652152687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417026595","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28169933,0.0030173957,0.70471054,0.00292279,0.00006496153,0.00024326482,0.00044106023,0.0017438531,0.005156758],"genre_scores_gemma":[0.8778983,0.0011416062,0.11568182,0.00036194795,0.00013675336,0.0004900671,0.000815784,0.0011370161,0.0023366776],"study_design_codex":"simulation_or_modeling","study_design_gemma":"observational","domain_scores_codex":[0.9496571,0.044985343,0.00082536385,0.0024034663,0.0016551723,0.0004734385],"domain_scores_gemma":[0.22624509,0.74220455,0.0071019,0.02062295,0.0030432688,0.00078231166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05642704,0.0017488376,0.0014692149,0.0029883548,0.0007910633,0.0031170866,0.0036012512,0.002415266,0.0047101784],"category_scores_gemma":[0.385297,0.0009453646,0.0016780372,0.0036754953,0.0029793398,0.011657476,0.0029857275,0.006062008,0.0013226924],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002477117,0.0019356097,0.049579706,0.0010637252,0.0017220274,0.00014895048,0.005335679,0.44873524,0.0028077431,0.15400125,0.008949274,0.32324362],"study_design_scores_gemma":[0.00031225142,0.00035416568,0.010207723,0.00016045506,0.00023286926,0.00013632704,0.00075303676,0.86896074,0.0018690195,0.11428139,0.0026446532,0.00008742771],"about_ca_topic_score_codex":0.005951645,"about_ca_topic_score_gemma":0.0045754514,"teacher_disagreement_score":0.05642704,"about_ca_system_score_codex":0.003247368,"about_ca_system_score_gemma":0.0021938027,"threshold_uncertainty_score":0.29841828},"labels":[],"label_agreement":null},{"id":"W4417130675","doi":"10.2139/ssrn.5887028","title":"MAC: Multi-Agent LLM Coder is All You Need","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Correctness; Robustness (evolution); Coding (social sciences); Code (set theory); Inference; Reliability (semiconductor)","score_opus":0.032895186743787365,"score_gpt":0.2928385054906594,"score_spread":0.259943318746872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417130675","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065631834,0.00045649978,0.6491752,0.005468188,0.0019823264,0.00032570073,0.0033169368,0.29199514,0.04071681],"genre_scores_gemma":[0.27598426,0.00068884116,0.526742,0.00622621,0.0017390756,0.0014467862,0.006152554,0.05129705,0.12972325],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99889684,0.0003041271,0.000085113796,0.00018001089,0.00037945798,0.00015436196],"domain_scores_gemma":[0.99462885,0.0013851185,0.00032718846,0.0015256383,0.001542186,0.0005910126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017956524,0.00094710273,0.00090198073,0.00073662336,0.00082643086,0.0023242515,0.001472885,0.0017398131,0.07936697],"category_scores_gemma":[0.01176845,0.0005661635,0.00051761797,0.00046087994,0.00067694613,0.0029413942,0.00230666,0.0030029367,0.052232027],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013006728,0.00021305506,0.0013508163,0.00041892816,0.00009260198,0.00034698783,0.00036678516,0.0123962965,0.012044949,0.047631666,0.5876141,0.3362231],"study_design_scores_gemma":[0.00047009595,0.00020422695,0.00074324006,0.00024214188,0.00006306602,0.00037532632,0.00012337488,0.3378702,0.027444094,0.0690342,0.563264,0.00016596472],"about_ca_topic_score_codex":0.0018324048,"about_ca_topic_score_gemma":0.0018903518,"teacher_disagreement_score":0.07936697,"about_ca_system_score_codex":0.00070147176,"about_ca_system_score_gemma":0.0012173773,"threshold_uncertainty_score":0.26550895},"labels":[],"label_agreement":null},{"id":"W4417312634","doi":"10.1142/9789819824755_0028","title":"Asking the Right Questions: Benchmarking Large Language Models in the Development of Clinical Consultation Templates","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Alberta","funders":"","keywords":"Benchmarking; Template; Prioritization; Pipeline (software); Comparability; Matching (statistics); Dependency (UML); Salient","score_opus":0.04134578372500843,"score_gpt":0.3577342401636168,"score_spread":0.3163884564386084,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417312634","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6449609,0.0010127684,0.31759712,0.0013351713,0.00017811672,0.0015915621,0.0030523401,0.018992426,0.011279629],"genre_scores_gemma":[0.7728449,0.00023702242,0.2194743,0.00036012032,0.000022120295,0.0006330957,0.004490967,0.00069936895,0.0012381198],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9930803,0.0048632305,0.00046775388,0.00070026045,0.00072203105,0.00016637499],"domain_scores_gemma":[0.95614403,0.03767073,0.0009837111,0.0024507144,0.002064433,0.000686483],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011298935,0.0011355237,0.0005451665,0.0011659219,0.00048612192,0.001973562,0.0017707978,0.0011896646,0.0028206755],"category_scores_gemma":[0.04998963,0.0005445238,0.00081916165,0.00082816713,0.0006155282,0.001832703,0.001888459,0.0012552226,0.0009919013],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015601474,0.0008404937,0.018545084,0.00067811797,0.0002512094,0.00032393014,0.0018344859,0.771679,0.004535817,0.005709082,0.0083633,0.18567932],"study_design_scores_gemma":[0.00016328809,0.00028295402,0.0011021164,0.000048789952,0.00005420917,0.000057229383,0.00029460553,0.98576766,0.00511448,0.0030293493,0.0040493323,0.00003592064],"about_ca_topic_score_codex":0.013056774,"about_ca_topic_score_gemma":0.016952459,"teacher_disagreement_score":0.013056774,"about_ca_system_score_codex":0.0026912752,"about_ca_system_score_gemma":0.0036751463,"threshold_uncertainty_score":0.059755206},"labels":[],"label_agreement":null},{"id":"W4417337812","doi":"10.1109/iccubea65967.2025.11283997","title":"Bias Mitigation in Nlp Models Using Bert","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Adversarial system; Language model; Pipeline (software); Benchmark (surveying); Sentence; Variety (cybernetics); Comprehension; Training set","score_opus":0.10557080173459263,"score_gpt":0.3055102447555013,"score_spread":0.19993944302090869,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417337812","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.114649884,0.00083564385,0.87118834,0.0013370084,0.00017452108,0.00024242522,0.00055737456,0.006384195,0.0046305535],"genre_scores_gemma":[0.8134978,0.00032141394,0.17598441,0.0008855451,0.00012598906,0.00036725478,0.0018219735,0.00068040704,0.006315203],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99870837,0.00061667763,0.000058079226,0.00030415092,0.0001925288,0.00012024258],"domain_scores_gemma":[0.9963302,0.0025482355,0.00021316262,0.0005437887,0.000263957,0.00010067284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032884032,0.0013277472,0.0007231578,0.00055788725,0.0005703092,0.0011521599,0.0016046131,0.001255045,0.002404178],"category_scores_gemma":[0.009580995,0.0005055534,0.0008718009,0.00039347282,0.0011473267,0.0022807494,0.002444847,0.0028150775,0.001258544],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039439744,0.00014043537,0.0036265901,0.00015831551,0.000096261596,0.00023116503,0.00034060326,0.866464,0.008982011,0.015959142,0.0058140876,0.097792916],"study_design_scores_gemma":[0.000011868274,0.000037698417,0.0001301379,0.000010570257,0.0000073266574,0.000026189644,0.00001821301,0.98897636,0.002391529,0.007448785,0.0009342152,0.0000072651355],"about_ca_topic_score_codex":0.002967159,"about_ca_topic_score_gemma":0.0045136306,"teacher_disagreement_score":0.0032884032,"about_ca_system_score_codex":0.0010248169,"about_ca_system_score_gemma":0.0011587732,"threshold_uncertainty_score":0.017390966},"labels":[],"label_agreement":null},{"id":"W4417457716","doi":"10.1177/21501319251404193","title":"Balancing Model Complexity and Clinical Deployability in Deep Learning for Sociodemographic Information Extraction","year":2025,"lang":"en","type":"article","venue":"Journal of Primary Care & Community Health","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"North York General Hospital; University Health Network; University of Toronto","funders":"Canadian Institutes of Health Research","keywords":"Documentation; Deep learning; Convolutional neural network; Information overload; Artificial neural network; Binary classification; Health care; Unstructured data; Medical record; Data extraction","score_opus":0.060778008221763734,"score_gpt":0.3738910076647382,"score_spread":0.3131129994429745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417457716","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7869869,0.0017697115,0.20168811,0.0026773466,0.00011864936,0.0003378931,0.001139925,0.0023948022,0.0028866408],"genre_scores_gemma":[0.94410616,0.000300922,0.0527249,0.00030121286,0.00003972248,0.00023962212,0.0014250798,0.000065421445,0.0007970378],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99782145,0.0010238251,0.00023024678,0.00048729224,0.00023307216,0.00020408754],"domain_scores_gemma":[0.99133146,0.006546385,0.00040965018,0.00076987065,0.0007901276,0.00015256114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008947396,0.0016258971,0.00090954156,0.0010245108,0.0006550249,0.0021410312,0.0014689434,0.001216957,0.0012764534],"category_scores_gemma":[0.02112826,0.00069299096,0.0009869768,0.0009289569,0.00076284143,0.0033467107,0.0019564247,0.0023037614,0.00047114154],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014259181,0.0004494576,0.062873684,0.00036574638,0.00067690003,0.00024598528,0.00050960976,0.63051414,0.0063875876,0.0031468077,0.0036313296,0.28977278],"study_design_scores_gemma":[0.000042556956,0.00014191485,0.0027346471,0.00004065161,0.000085080515,0.000035602563,0.00008144535,0.9898274,0.00245859,0.0040234965,0.00051265466,0.000015979433],"about_ca_topic_score_codex":0.016522784,"about_ca_topic_score_gemma":0.019914003,"teacher_disagreement_score":0.016522784,"about_ca_system_score_codex":0.002026311,"about_ca_system_score_gemma":0.0023728441,"threshold_uncertainty_score":0.047318935},"labels":[],"label_agreement":null},{"id":"W4417462069","doi":"10.2139/ssrn.5735187","title":"Mitigating Lost-in-the-Middle: Dynamic Semantic Chunking for Precise Token-Level Retrieval in RAG","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Chunking (psychology); Security token; Coherence (philosophical gambling strategy); Embedding; Merge (version control); Relevance (law); Semantic role labeling","score_opus":0.03661898335840118,"score_gpt":0.28835293495432923,"score_spread":0.25173395159592804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417462069","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026705084,0.0010924838,0.95666707,0.0004370383,0.00027218537,0.00016255444,0.0006294981,0.011906162,0.0021278847],"genre_scores_gemma":[0.49447626,0.000539417,0.49618062,0.00038822266,0.0002829834,0.00021325759,0.0018827942,0.0016253943,0.0044110324],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99745494,0.00083166786,0.00027227227,0.00061225257,0.00045257117,0.00037632708],"domain_scores_gemma":[0.9922691,0.0029581457,0.0003026216,0.0034873232,0.0006974759,0.00028548727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004062415,0.0010580367,0.0020235337,0.001875431,0.00145867,0.0029844588,0.0032859708,0.0022523263,0.0075949645],"category_scores_gemma":[0.014809286,0.0007798352,0.0008885988,0.0022231848,0.001832499,0.0094802715,0.0061975108,0.0027181362,0.0049645226],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017404658,0.0003253616,0.0024360127,0.0006096944,0.00018577454,0.0004062562,0.0024187607,0.048047807,0.052628774,0.047102373,0.025521636,0.8185771],"study_design_scores_gemma":[0.00012571418,0.00028728181,0.0008324513,0.00008165265,0.00017702456,0.0003091357,0.0008925334,0.7710479,0.058452107,0.1471431,0.020516811,0.00013422442],"about_ca_topic_score_codex":0.0041895583,"about_ca_topic_score_gemma":0.005098674,"teacher_disagreement_score":0.0075949645,"about_ca_system_score_codex":0.0008815793,"about_ca_system_score_gemma":0.0026820248,"threshold_uncertainty_score":0.025407732},"labels":[],"label_agreement":null},{"id":"W4417469644","doi":"10.48550/arxiv.2512.14427","title":"Effect of Document Packing on the Latent Multi-Hop Reasoning Capabilities of Large Language Models","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alliance de recherche numérique du Canada; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Key (lock); Process (computing); Language model; Language understanding; Data modeling; Training set","score_opus":0.035332971124111275,"score_gpt":0.30049031727439146,"score_spread":0.2651573461502802,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417469644","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92746186,0.001456967,0.063795485,0.0014342954,0.00012766873,0.000101295416,0.00029965505,0.0013844841,0.0039382176],"genre_scores_gemma":[0.9743251,0.00035176345,0.023786087,0.0002276212,0.000048469316,0.000069759146,0.00037774572,0.00020550903,0.0006080415],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9962274,0.0021264441,0.00028638952,0.0005614468,0.00037517794,0.0004231392],"domain_scores_gemma":[0.84351534,0.14070982,0.0030751834,0.008629456,0.0020479835,0.0020223004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0104341945,0.0011680311,0.0012622314,0.000911964,0.0013211714,0.0024818457,0.0013956717,0.0020181637,0.0026959905],"category_scores_gemma":[0.11152908,0.000989815,0.0008301337,0.0009885578,0.002210099,0.006896983,0.00294066,0.0037681658,0.0006341182],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005315095,0.001237525,0.027368125,0.0007242005,0.0003643421,0.00053263525,0.0020167036,0.70223,0.017610367,0.01340919,0.005565889,0.22362597],"study_design_scores_gemma":[0.00017683824,0.00070182356,0.003874552,0.0001018144,0.00014120816,0.00018822287,0.00050932134,0.9629323,0.012769522,0.017514078,0.0010378128,0.000052520085],"about_ca_topic_score_codex":0.0043442254,"about_ca_topic_score_gemma":0.004393838,"teacher_disagreement_score":0.0104341945,"about_ca_system_score_codex":0.0012338244,"about_ca_system_score_gemma":0.0017474971,"threshold_uncertainty_score":0.05518198},"labels":[],"label_agreement":null},{"id":"W4417509732","doi":"10.1109/aibthings66987.2025.11296175","title":"From Large to Small Language Models for Balanced Performance and Benchmarking","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Benchmarking; Inference; Baseline (sea); Sustainable development; Language model; Natural language understanding; Natural language; Comprehension","score_opus":0.019264700670818948,"score_gpt":0.25469516079173876,"score_spread":0.2354304601209198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417509732","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44548362,0.005993187,0.39361784,0.005519398,0.0016100215,0.0009712506,0.01146313,0.084810674,0.05053098],"genre_scores_gemma":[0.70522696,0.0009507161,0.26984003,0.0008513607,0.00009846492,0.00074680516,0.013552609,0.004282507,0.0044504534],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99661225,0.0013051503,0.0002604959,0.000725179,0.0008213421,0.0002755395],"domain_scores_gemma":[0.9946063,0.002613925,0.00015058821,0.0015114495,0.0008418277,0.00027596005],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038172056,0.0014834752,0.0008733206,0.0010395116,0.0007062619,0.00237975,0.002748495,0.0014207514,0.005675667],"category_scores_gemma":[0.018845107,0.00059311936,0.0010190017,0.001408802,0.0008063106,0.0048619187,0.0020588161,0.0025158955,0.0030826295],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019676997,0.0011629397,0.011907807,0.0020340874,0.00076042884,0.00046636496,0.0009587923,0.5122013,0.02166474,0.03272664,0.095104344,0.31904483],"study_design_scores_gemma":[0.00016986468,0.00024596634,0.0010115533,0.00008444343,0.00007503452,0.00007474406,0.00030905037,0.9440835,0.01435307,0.020141928,0.01940825,0.000042471354],"about_ca_topic_score_codex":0.013240528,"about_ca_topic_score_gemma":0.018143263,"teacher_disagreement_score":0.013240528,"about_ca_system_score_codex":0.0016073458,"about_ca_system_score_gemma":0.0026011274,"threshold_uncertainty_score":0.026326895},"labels":[],"label_agreement":null},{"id":"W4417510529","doi":"10.1109/eecsi67060.2025.11290439","title":"Faithful by Design: Improving Large Language Model Rationales through Counterfactual Consistency Verification","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Marriott International (Canada)","funders":"","keywords":"Counterfactual thinking; Consistency (knowledge bases); Benchmark (surveying); Generalizability theory; Language model; Natural language understanding; Limiting; Inference; Ranging","score_opus":0.029943144516298325,"score_gpt":0.28022198745671945,"score_spread":0.25027884294042113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417510529","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02737835,0.0003563714,0.9557527,0.0012213767,0.00008880058,0.00038602893,0.000575518,0.012413788,0.0018271377],"genre_scores_gemma":[0.29118532,0.00020465684,0.70260495,0.0005057549,0.000056124878,0.00036777786,0.002246324,0.001748183,0.001080785],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9869376,0.0070615313,0.0007644892,0.0015046122,0.0033080007,0.0004237713],"domain_scores_gemma":[0.9457532,0.033486597,0.0029678268,0.013563201,0.0037413498,0.00048782813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018126778,0.0016133753,0.0010493601,0.0024104351,0.0007418628,0.004210353,0.0038263625,0.0020095864,0.0046707164],"category_scores_gemma":[0.095393464,0.0009713429,0.0032587673,0.0008650487,0.0027048988,0.0060796947,0.0046625375,0.0032062896,0.0012944968],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092625764,0.0006117479,0.01707438,0.0014353446,0.00045942282,0.0011611105,0.0024268471,0.278203,0.026492054,0.13750865,0.01283444,0.52086675],"study_design_scores_gemma":[0.00014564773,0.00015455521,0.000621436,0.00016699902,0.00012696645,0.00023061427,0.00017782887,0.87849534,0.012533757,0.09854202,0.008748104,0.000056731544],"about_ca_topic_score_codex":0.0037195892,"about_ca_topic_score_gemma":0.0062394524,"teacher_disagreement_score":0.018126778,"about_ca_system_score_codex":0.0016948389,"about_ca_system_score_gemma":0.0048159054,"threshold_uncertainty_score":0.09586465},"labels":[],"label_agreement":null},{"id":"W49335612","doi":"10.1007/978-3-642-40769-7_7","title":"R/quest: A Question Answering System","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Computer science; Question answering; Cosine similarity; Information retrieval; Ranking (information retrieval); Similarity (geometry); Rank (graph theory); Mean reciprocal rank; Search engine; Domain (mathematical analysis); Recall; Open domain; Precision and recall; Artificial intelligence; Mathematics; Pattern recognition (psychology)","score_opus":0.016187404134827568,"score_gpt":0.23397914192795982,"score_spread":0.21779173779313227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W49335612","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00839606,0.0012122128,0.66759336,0.0012334861,0.00033721313,0.00049539976,0.019874593,0.2522498,0.048607927],"genre_scores_gemma":[0.06829916,0.0010444211,0.78834534,0.0016053225,0.00031992048,0.000761607,0.05514387,0.018821398,0.06565889],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9992887,0.00017311184,0.00007670086,0.00023280797,0.00018886711,0.000039864128],"domain_scores_gemma":[0.9988763,0.00066810113,0.000040037754,0.00017672732,0.000162024,0.00007682129],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015714661,0.0012197347,0.0011912113,0.002027584,0.0006893037,0.0023793243,0.0018359342,0.0013526651,0.04122919],"category_scores_gemma":[0.0038367237,0.00097609556,0.0012151101,0.0013460041,0.00062697974,0.0050520035,0.0026817801,0.0013818284,0.04219304],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060778874,0.0001684876,0.0010855577,0.0010869103,0.00006978314,0.00029082244,0.0006762057,0.0025135542,0.020095252,0.056236554,0.4760019,0.44116718],"study_design_scores_gemma":[0.0003666291,0.00019432273,0.0010631431,0.00023085593,0.00012797424,0.0009093277,0.00036248082,0.121578,0.026698003,0.11985185,0.72848105,0.00013639167],"about_ca_topic_score_codex":0.0014652293,"about_ca_topic_score_gemma":0.0018965761,"teacher_disagreement_score":0.04122919,"about_ca_system_score_codex":0.0005900594,"about_ca_system_score_gemma":0.00075716304,"threshold_uncertainty_score":0.13792539},"labels":[],"label_agreement":null},{"id":"W53236038","doi":"10.1007/978-3-319-07983-7_21","title":"Complex Question Answering: Homogeneous or Heterogeneous, Which Ensemble Is Better?","year":2014,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Homogeneous; Support vector machine; Hidden Markov model; Ensemble learning; Conditional random field; Artificial intelligence; Maximum-entropy Markov model; Question answering; Machine learning; Base (topology); Entropy (arrow of time); Principle of maximum entropy; Ensemble forecasting; Markov chain; Task (project management); Markov model; Variable-order Markov model; Mathematics","score_opus":0.030604614856074654,"score_gpt":0.26162371290475506,"score_spread":0.2310190980486804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W53236038","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14413258,0.018597648,0.79048735,0.011582833,0.0008098147,0.00030226677,0.002473606,0.0032979792,0.028315984],"genre_scores_gemma":[0.6559421,0.004975745,0.31591743,0.0025137982,0.0036023501,0.00025253158,0.0070835417,0.0010154827,0.008697074],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960085,0.0013151135,0.0002375728,0.0015069642,0.0006992523,0.00023260144],"domain_scores_gemma":[0.9865536,0.007931422,0.00058538886,0.0025756902,0.001526388,0.0008274609],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077853054,0.001137403,0.0020190976,0.0018605943,0.000803634,0.0042579356,0.0018919557,0.0018025065,0.007725801],"category_scores_gemma":[0.018174171,0.00054649706,0.0015715499,0.0021472843,0.001089579,0.010814836,0.003088043,0.0027279102,0.0036470229],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008310488,0.00050974754,0.016935539,0.001072249,0.0011494253,0.00011035025,0.0011017785,0.012146455,0.01481227,0.025696794,0.039855998,0.88577825],"study_design_scores_gemma":[0.0001940008,0.0006876354,0.03250322,0.0005044957,0.0017639041,0.0013529741,0.0016265127,0.4201134,0.02205429,0.45052502,0.068469465,0.00020498481],"about_ca_topic_score_codex":0.00069171813,"about_ca_topic_score_gemma":0.0010680885,"teacher_disagreement_score":0.0077853054,"about_ca_system_score_codex":0.0005992728,"about_ca_system_score_gemma":0.0005877707,"threshold_uncertainty_score":0.0411731},"labels":[],"label_agreement":null},{"id":"W54573516","doi":"","title":"Dialogical Models of Explanation.","year":2007,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg","funders":"","keywords":"Dialogical self; Epistemology; Proposition; Set (abstract data type); Computer science; Inference; Simple (philosophy); Task (project management); Closing (real estate); Artificial intelligence; Cognitive science; Psychology; Philosophy; Law; Political science","score_opus":0.06199680857406427,"score_gpt":0.2655703515198456,"score_spread":0.20357354294578134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W54573516","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005729655,0.022573296,0.4877986,0.023583207,0.0010383943,0.00022162532,0.00091410906,0.0008851036,0.45725602],"genre_scores_gemma":[0.6713384,0.011938603,0.20924033,0.0050099073,0.0022118443,0.00091587123,0.0017621199,0.0004051836,0.09717767],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99643904,0.0018572157,0.00023134248,0.00068664545,0.000557009,0.00022874257],"domain_scores_gemma":[0.9961074,0.0023712704,0.0002956385,0.00071036286,0.00030811573,0.00020718171],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036596567,0.0010164077,0.00068996847,0.001863957,0.0023361272,0.005901254,0.0018564439,0.0040997346,0.027444843],"category_scores_gemma":[0.0053211395,0.0006216665,0.0016446109,0.0012426806,0.008601248,0.011529704,0.0038993927,0.003223584,0.0033693626],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000004163414,0.0000037261314,0.000046323104,0.000033421777,0.0000050358435,0.000026484051,0.00030996156,0.00045209684,0.000027607846,0.9952689,0.0010262172,0.0027960625],"study_design_scores_gemma":[0.000010445961,0.000007732486,0.0000557444,0.00004560616,0.000006480239,0.00008479974,0.00014546573,0.0029834001,0.00006304397,0.94626147,0.050328813,0.0000068334416],"about_ca_topic_score_codex":0.003194792,"about_ca_topic_score_gemma":0.0022649674,"teacher_disagreement_score":0.027444843,"about_ca_system_score_codex":0.0036771162,"about_ca_system_score_gemma":0.0022386985,"threshold_uncertainty_score":0.091812134},"labels":[],"label_agreement":null},{"id":"W57814663","doi":"10.1007/3-540-36618-0_24","title":"Combining Naive Bayes and n-Gram Language Models for Text Classification","year":2003,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":170,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Naive Bayes classifier; Computer science; Artificial intelligence; Machine learning; Bayes error rate; Bayesian programming; Bayes' theorem; n-gram; Smoothing; Conditional independence; Inference; Bayes classifier; Bayes factor; Language model; Bayesian probability; Support vector machine","score_opus":0.03324207233675179,"score_gpt":0.26568466222777093,"score_spread":0.23244258989101912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W57814663","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01138954,0.00264891,0.97767305,0.00064806367,0.0005076487,0.0002083241,0.0006260628,0.004469686,0.0018286179],"genre_scores_gemma":[0.19662988,0.0027266976,0.7854684,0.00085179886,0.0016013584,0.0005247526,0.0040879794,0.0006410909,0.007468045],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99320847,0.0031070928,0.0005814174,0.0011405464,0.0016585446,0.0003039197],"domain_scores_gemma":[0.98703057,0.009381073,0.0003173358,0.0008626188,0.0021602719,0.00024810672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0079470305,0.0022606526,0.003344682,0.0049699214,0.0018727677,0.0032679213,0.0026519015,0.0024028965,0.0036318863],"category_scores_gemma":[0.018329363,0.0010915634,0.0024659622,0.0043020155,0.0007997833,0.007657116,0.0016420323,0.003156049,0.006397765],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065618946,0.00050292193,0.0026445203,0.00046331406,0.00040703025,0.00014906185,0.00019159408,0.028592508,0.006200324,0.0053376346,0.02014458,0.93471026],"study_design_scores_gemma":[0.0000660967,0.000111620175,0.0007435347,0.000085208594,0.00025867074,0.0001948767,0.00009106134,0.9457187,0.0041423845,0.044580825,0.0039360244,0.00007100831],"about_ca_topic_score_codex":0.010003686,"about_ca_topic_score_gemma":0.017693343,"teacher_disagreement_score":0.010003686,"about_ca_system_score_codex":0.0011092161,"about_ca_system_score_gemma":0.0023305514,"threshold_uncertainty_score":0.042028427},"labels":[],"label_agreement":null},{"id":"W584835505","doi":"","title":"Automatic text summarization in digital libraries","year":2006,"lang":"en","type":"dissertation","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Digital library; Computer science; Information retrieval; World Wide Web; Natural language processing; Library science; Linguistics; Philosophy","score_opus":0.00887390532647554,"score_gpt":0.21874313448194246,"score_spread":0.20986922915546694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W584835505","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13376163,0.004787139,0.8176876,0.001122233,0.00029122544,0.00087558187,0.0024335918,0.026394159,0.012646789],"genre_scores_gemma":[0.30680054,0.0025833678,0.66674757,0.00018905026,0.00043352207,0.00062580075,0.008126102,0.0014158428,0.01307815],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99805164,0.00070603815,0.00020565132,0.00038166074,0.0005544757,0.00010058668],"domain_scores_gemma":[0.99505204,0.0025485659,0.0005323344,0.00037049173,0.0013800719,0.00011647433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016291947,0.00082504255,0.00093218515,0.004700803,0.0008708738,0.0026579178,0.0010417382,0.00067406846,0.0045518987],"category_scores_gemma":[0.009223356,0.00048581968,0.0006861884,0.0043663145,0.0003512616,0.0029237347,0.00102648,0.00055083184,0.0037773848],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003768831,0.000114889706,0.0017468963,0.00090071425,0.00006208507,0.00017786033,0.0019018904,0.009479273,0.025999537,0.0027925922,0.013423305,0.9430242],"study_design_scores_gemma":[0.00036194533,0.0013079926,0.020660486,0.0005088335,0.0007087627,0.0011451441,0.008592464,0.5364279,0.20011312,0.02885224,0.20104136,0.0002797828],"about_ca_topic_score_codex":0.0025479174,"about_ca_topic_score_gemma":0.0024541016,"teacher_disagreement_score":0.004700803,"about_ca_system_score_codex":0.000761039,"about_ca_system_score_gemma":0.0007780085,"threshold_uncertainty_score":0.015227556},"labels":[],"label_agreement":null},{"id":"W59336092","doi":"10.1007/978-3-642-30353-1_39","title":"Exploiting Semantic Roles for Asynchronous Question Answering in an Educational Setting","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Athabasca University","funders":"","keywords":"Asynchronous communication; Question answering; Computer science; Asynchronous learning; Information retrieval; Natural language processing; Psychology; Mathematics education; Teaching method","score_opus":0.021563959195854278,"score_gpt":0.2759686408468404,"score_spread":0.2544046816509861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W59336092","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042691864,0.00028607823,0.9472536,0.0008510021,0.00011815219,0.00012356893,0.00033072452,0.002813909,0.0055311085],"genre_scores_gemma":[0.6333496,0.0002120008,0.3613541,0.00017519089,0.00012511412,0.00012910893,0.0010577803,0.00031069692,0.0032864197],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99638546,0.0020825197,0.00023050181,0.00062441814,0.00041519853,0.0002617753],"domain_scores_gemma":[0.98921824,0.007945416,0.00039364686,0.0012298936,0.00071942265,0.0004934032],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045080227,0.00080966984,0.00067393866,0.001079101,0.0010299492,0.003015666,0.0020752896,0.0017334396,0.0057395287],"category_scores_gemma":[0.014254042,0.00074590347,0.001186115,0.0009226627,0.0010091653,0.008909556,0.0034675489,0.0022443305,0.0016954367],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024975701,0.0013853136,0.008590438,0.0010146606,0.00016545638,0.00123273,0.008992563,0.036408763,0.06635445,0.32685903,0.01798562,0.5285133],"study_design_scores_gemma":[0.0001869225,0.00019101829,0.0012986029,0.00006847583,0.00015222395,0.0003284912,0.0013952451,0.61449754,0.021268671,0.33547455,0.025061741,0.000076457865],"about_ca_topic_score_codex":0.0017739454,"about_ca_topic_score_gemma":0.0025745647,"teacher_disagreement_score":0.0057395287,"about_ca_system_score_codex":0.0007342807,"about_ca_system_score_gemma":0.0008603455,"threshold_uncertainty_score":0.023840964},"labels":[],"label_agreement":null},{"id":"W60608539","doi":"","title":"Passage Retrieval by Shrinkage of Language Models.","year":2006,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Computer science; Information retrieval; Scope (computer science); Search engine; World Wide Web; The Internet; Process (computing); Human–computer information retrieval","score_opus":0.02283974409388151,"score_gpt":0.2468398962509053,"score_spread":0.2240001521570238,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W60608539","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00502526,0.002362251,0.9868326,0.00045305802,0.00033531204,0.00021103422,0.0005454954,0.0016922068,0.0025427986],"genre_scores_gemma":[0.2612721,0.0071674017,0.6777792,0.0011566331,0.002324821,0.0022741489,0.0074949544,0.0023464411,0.038184308],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960006,0.0021824064,0.00021631761,0.0007439905,0.0006266758,0.00023001165],"domain_scores_gemma":[0.98751324,0.009126886,0.00054205954,0.0015580406,0.0010059774,0.00025375962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0093001705,0.0018604704,0.0024990244,0.0034483317,0.0009514055,0.0024076402,0.002700091,0.0018553956,0.008803787],"category_scores_gemma":[0.03729002,0.001320247,0.003052324,0.0029266193,0.0016688402,0.005749535,0.0027975633,0.003694678,0.008370481],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000777949,0.00033658298,0.0020211714,0.0009978095,0.00056372257,0.00050028425,0.0007464394,0.27592093,0.006187542,0.24987532,0.0380997,0.42397264],"study_design_scores_gemma":[0.0000723843,0.00011464274,0.00043996953,0.00005797728,0.00012877145,0.00021175339,0.000043867472,0.8614446,0.0019811692,0.12348616,0.011959539,0.000059049788],"about_ca_topic_score_codex":0.0060057305,"about_ca_topic_score_gemma":0.006388693,"teacher_disagreement_score":0.0093001705,"about_ca_system_score_codex":0.0014604882,"about_ca_system_score_gemma":0.0015913282,"threshold_uncertainty_score":0.04918462},"labels":[],"label_agreement":null},{"id":"W6261272","doi":"10.1007/978-3-642-21043-3_9","title":"Using Semantic Information to Answer Complex Questions","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science; Semantic similarity; Artificial intelligence; Natural language processing; Word (group theory); Graph; String (physics); Benchmark (surveying); Question answering; Subsequence; Information retrieval; Theoretical computer science; Mathematics","score_opus":0.05982960648080021,"score_gpt":0.2787075651891551,"score_spread":0.2188779587083549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6261272","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0102731,0.0026498286,0.9442523,0.0024310513,0.00048771861,0.00019343723,0.0013966425,0.0033250374,0.034990855],"genre_scores_gemma":[0.14179689,0.0034445752,0.8239886,0.0007400149,0.0005119616,0.00035965786,0.0059506255,0.00070231385,0.022505356],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99884534,0.0004037522,0.0000877482,0.00024340911,0.00036229877,0.000057349167],"domain_scores_gemma":[0.9968721,0.002452043,0.0001043104,0.00025272398,0.00025013747,0.00006865947],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015368039,0.0010115544,0.0006329414,0.0019537923,0.00066181825,0.0025445614,0.001217186,0.0012502148,0.012716193],"category_scores_gemma":[0.00760578,0.0004812566,0.0013603625,0.0015095292,0.0011249289,0.007633135,0.0017496479,0.0020066681,0.004893075],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019952875,0.00015649531,0.0009284069,0.001361237,0.000115056915,0.0001748874,0.0023230324,0.006530785,0.014884982,0.27902782,0.048156414,0.64614135],"study_design_scores_gemma":[0.000063531275,0.000078042394,0.0008146836,0.00032884034,0.0001524994,0.00026111116,0.0007361012,0.079814136,0.010888181,0.7053376,0.20146878,0.000056613873],"about_ca_topic_score_codex":0.0010091845,"about_ca_topic_score_gemma":0.0015603438,"teacher_disagreement_score":0.012716193,"about_ca_system_score_codex":0.00062388374,"about_ca_system_score_gemma":0.0007269443,"threshold_uncertainty_score":0.042539954},"labels":[],"label_agreement":null},{"id":"W63277956","doi":"10.1007/978-3-642-30353-1_33","title":"A Three-Level Cognitive Architecture for the Simulation of Human Behaviour","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Cognitive architecture; Computer science; Architecture; Cognition; Human–computer interaction; Cognitive science; Stroop effect; Task (project management); Cognitive model; Artificial intelligence; Psychology; Systems engineering","score_opus":0.06882153734754258,"score_gpt":0.3029974186831052,"score_spread":0.2341758813355626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W63277956","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03016426,0.00017171087,0.9530956,0.0005382663,0.000045743804,0.00009492801,0.0001775318,0.0015677072,0.0141442185],"genre_scores_gemma":[0.47907522,0.00030059612,0.5143649,0.00012995118,0.000022722244,0.0003169289,0.00025255946,0.00015689581,0.0053802123],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99970526,0.0001101816,0.00002187086,0.000059863523,0.00006746381,0.0000352736],"domain_scores_gemma":[0.9994332,0.00028044256,0.000030195246,0.00010932287,0.000077306104,0.00006948706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006068722,0.00051512744,0.00048441207,0.00047924859,0.0009731389,0.0029875804,0.002134501,0.0015465638,0.007460431],"category_scores_gemma":[0.0022283832,0.00047619647,0.0013843094,0.00046987052,0.0016843468,0.002569626,0.0015310184,0.001501075,0.0010257902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015588255,0.00012616701,0.0017609204,0.00021109784,0.00012225633,0.00024750005,0.002211409,0.40766066,0.008786551,0.53121674,0.0024541677,0.045046706],"study_design_scores_gemma":[0.000025530679,0.00004255087,0.00036337023,0.000028107248,0.00003776934,0.000055915745,0.00009908311,0.75142336,0.0012512034,0.24124914,0.00540078,0.000023116554],"about_ca_topic_score_codex":0.010860771,"about_ca_topic_score_gemma":0.0076052574,"teacher_disagreement_score":0.010860771,"about_ca_system_score_codex":0.0011655665,"about_ca_system_score_gemma":0.001496426,"threshold_uncertainty_score":0.024957657},"labels":[],"label_agreement":null},{"id":"W65714572","doi":"","title":"Wikispeedia: an online game for inferring semantic distances between concepts","year":2009,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Latent semantic analysis; Point (geometry); Semantics (computer science); Artificial intelligence; Semantic similarity; Natural language processing; Semantic computing; Quality (philosophy); Information retrieval; Semantic Web; Data science","score_opus":0.07418840991691421,"score_gpt":0.343737642741528,"score_spread":0.2695492328246138,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W65714572","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.120478965,0.0007263035,0.8445558,0.0003073508,0.00022060616,0.000939742,0.013687003,0.013657306,0.0054268935],"genre_scores_gemma":[0.25207055,0.00027814202,0.7277152,0.00009496317,0.00004988779,0.0007940409,0.016153978,0.0003401977,0.002503066],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991561,0.00030775435,0.00006598779,0.0002496586,0.00018030968,0.00004016165],"domain_scores_gemma":[0.9978849,0.001289088,0.00017493559,0.00027159046,0.00023633821,0.00014302737],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00092807226,0.0016327415,0.00053510687,0.003348472,0.00058283086,0.0013608924,0.0020943654,0.00088318763,0.0022036866],"category_scores_gemma":[0.005939926,0.0004589558,0.00076234806,0.0013842223,0.00044809552,0.0033166,0.0020291305,0.0012832127,0.0012372461],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001505465,0.0014734716,0.03493278,0.0015159217,0.000754466,0.0007437639,0.0030441657,0.029428544,0.03328318,0.033438895,0.045193005,0.8146863],"study_design_scores_gemma":[0.00016113186,0.0004696891,0.020110918,0.00012169275,0.0001616111,0.00089843164,0.000901059,0.8548596,0.02123656,0.042566847,0.058331594,0.00018089179],"about_ca_topic_score_codex":0.0058917333,"about_ca_topic_score_gemma":0.016372802,"teacher_disagreement_score":0.0058917333,"about_ca_system_score_codex":0.00063811895,"about_ca_system_score_gemma":0.0008642454,"threshold_uncertainty_score":0.011714876},"labels":[],"label_agreement":null},{"id":"W65944014","doi":"10.1609/icwsm.v3i1.13980","title":"Regression-Based Summarization of Email Conversations","year":2009,"lang":"en","type":"article","venue":"Proceedings of the International AAAI Conference on Web and Social Media","topic":"Topic Modeling","field":"Computer Science","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Automatic summarization; Computer science; Artificial intelligence; Regression; Binary classification; Natural language processing; Machine learning; Random forest; Regression analysis; Recall; Support vector machine; Statistics; Mathematics","score_opus":0.03218593799524793,"score_gpt":0.25520677507144235,"score_spread":0.2230208370761944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W65944014","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07558882,0.0016850209,0.9027947,0.00050245546,0.0003444437,0.0003831105,0.0024626153,0.013314399,0.0029244144],"genre_scores_gemma":[0.47273523,0.00086524663,0.50424176,0.00016421375,0.00072501856,0.0005242136,0.0130030755,0.0011441376,0.0065971394],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966543,0.0014134402,0.00026328547,0.00081684213,0.0006736053,0.00017846398],"domain_scores_gemma":[0.9890103,0.005029781,0.0013989846,0.0009776346,0.003367493,0.0002158652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031417815,0.0015292089,0.0013613177,0.0043546227,0.0006408338,0.0015025502,0.0011796296,0.0009477834,0.0022181694],"category_scores_gemma":[0.017659876,0.00037656698,0.0008715882,0.0023618517,0.00026530187,0.0024330656,0.00084889086,0.0015119778,0.0032718629],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006744264,0.00026675418,0.0069548944,0.0007374614,0.00022924462,0.00016291557,0.00088924926,0.04072216,0.04092926,0.0023968779,0.012378857,0.8936579],"study_design_scores_gemma":[0.000028340386,0.00030796032,0.011665931,0.00008983206,0.00016430314,0.00016706635,0.00042613092,0.9384579,0.029504254,0.0066424976,0.01246284,0.00008292181],"about_ca_topic_score_codex":0.0019514387,"about_ca_topic_score_gemma":0.0023351791,"teacher_disagreement_score":0.0043546227,"about_ca_system_score_codex":0.0006476454,"about_ca_system_score_gemma":0.00058776786,"threshold_uncertainty_score":0.01661551},"labels":[],"label_agreement":null},{"id":"W68132019","doi":"10.1007/s10994-013-5363-6","title":"A semantic matching energy function for learning with multi-relational data","year":2013,"lang":"en","type":"article","venue":"Machine Learning","topic":"Topic Modeling","field":"Computer Science","cited_by":687,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Agence Nationale de la Recherche; Compute Canada; Defense Advanced Research Projects Agency; Canadian Institute for Advanced Research","keywords":"Computer science; WordNet; Artificial intelligence; Statistical relational learning; Natural language processing; Sentence; Relational database; Relation (database); Semantic matching; Semantics (computer science); Parsing; Word (group theory); Context (archaeology); Information retrieval; Matching (statistics); Data mining","score_opus":0.03555843756672505,"score_gpt":0.24708839161885043,"score_spread":0.21152995405212538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W68132019","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060327277,0.0003473289,0.9916203,0.00027482503,0.00006110873,0.00005359745,0.00017758137,0.0003534572,0.0010790235],"genre_scores_gemma":[0.256939,0.00078970235,0.7314159,0.00041199662,0.0002138751,0.0005567508,0.001829553,0.0007069246,0.0071363025],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99821424,0.0005729099,0.00016400045,0.00035290676,0.0005693087,0.00012674244],"domain_scores_gemma":[0.9982937,0.000853479,0.00008196442,0.00042262586,0.0002750463,0.00007316673],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034982867,0.0005632806,0.0014756514,0.002029199,0.0007871725,0.0019347112,0.0025056242,0.002531923,0.005321831],"category_scores_gemma":[0.009402822,0.00049657456,0.0016091021,0.0032342058,0.00088537496,0.0050688186,0.0028691085,0.002297066,0.0015143495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028063057,0.0004413804,0.0017621187,0.00037990883,0.000263755,0.00016921412,0.00023550414,0.20191793,0.0076524517,0.23260112,0.013289945,0.541006],"study_design_scores_gemma":[0.00001461601,0.000047440673,0.00041003738,0.00003027364,0.0000341292,0.00009620668,0.000044760447,0.8326504,0.0014688442,0.16193864,0.003243565,0.000021069123],"about_ca_topic_score_codex":0.0015199851,"about_ca_topic_score_gemma":0.0018990347,"teacher_disagreement_score":0.005321831,"about_ca_system_score_codex":0.001106068,"about_ca_system_score_gemma":0.0010790655,"threshold_uncertainty_score":0.018500924},"labels":[],"label_agreement":null},{"id":"W6889294309","doi":"10.25549/pcra-c14-193952","title":"Khristiaskiy vestnik = [Christian herald], vol. 25, whole no. 268 (1961 March)","year":2021,"lang":"en","type":"dataset","venue":"University of Southern California Digital Library","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Slavic languages; Diaspora; Faith; China; Slavic studies; Christianity","score_opus":0.0077660303307138905,"score_gpt":0.1646396128503074,"score_spread":0.15687358251959352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6889294309","genre_codex":"other","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0023066944,0.2644931,0.00057017506,0.0117146345,0.027321119,0.000066389315,0.0015927982,0.00029907707,0.691636],"genre_scores_gemma":[0.031372994,0.093509644,0.00040161394,0.002668381,0.006779944,0.00008699334,0.000879071,0.0003899273,0.86391145],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996737,0.000056441222,0.000032799842,0.00006897257,0.00012103493,0.000047007226],"domain_scores_gemma":[0.9997693,0.00006831122,0.00003852788,0.000025242321,0.00006915911,0.000029566678],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00049027684,0.0010994053,0.00045091414,0.001793051,0.0024378817,0.0037745778,0.0004993013,0.0010510749,0.09570611],"category_scores_gemma":[0.0013969077,0.0004931587,0.0002208966,0.0027341505,0.0011958834,0.0028138838,0.0015107256,0.0018051192,0.046211105],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000068255715,0.000015639876,0.0003294481,0.00041616446,0.000006614293,0.000092127244,0.00053384487,0.00007103864,0.0002837535,0.02845969,0.9055342,0.064189315],"study_design_scores_gemma":[0.0000033600206,0.00000413035,0.00058925146,0.00010641927,0.0000012658878,0.00008859515,0.00011622014,0.000008500232,0.0000455102,0.0008978018,0.9981363,0.0000027344613],"about_ca_topic_score_codex":0.00815102,"about_ca_topic_score_gemma":0.012765886,"teacher_disagreement_score":0.09570611,"about_ca_system_score_codex":0.0016888919,"about_ca_system_score_gemma":0.0011661558,"threshold_uncertainty_score":0.3201688},"labels":[],"label_agreement":null},{"id":"W6890032672","doi":"10.3389/fvets.2020.00388.s002","title":"Data_Sheet_2_Assisting Decision-Making on Age of Neutering for 35 Breeds of Dogs: Associated Joint Disorders, Cancers, and Urinary Incontinence.pdf","year":2020,"lang":"en","type":"dataset","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Neutering; Stifle joint; German Shepherd Dog; Breed; Labrador Retriever; Cruciate ligament; Epidemiology","score_opus":0.04614033846536296,"score_gpt":0.2860301656460406,"score_spread":0.23988982718067764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6890032672","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025368268,0.000113779664,0.0007406789,0.0009768145,0.00010363581,0.0019557346,0.9774041,0.00085535273,0.015313073],"genre_scores_gemma":[0.032651354,0.0009692355,0.013432028,0.0026514272,0.00027486274,0.020471485,0.9041334,0.0005570211,0.02485925],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985606,0.00037804423,0.0004888746,0.00013599898,0.00029652467,0.00013999741],"domain_scores_gemma":[0.97182435,0.01274914,0.004014031,0.0013817268,0.008745321,0.001285436],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0033237499,0.00074214785,0.000935742,0.0035067077,0.00080359343,0.0014975879,0.0015258371,0.0012001118,0.27151084],"category_scores_gemma":[0.02272382,0.0006375475,0.00086441386,0.0034882363,0.00023478724,0.0015079862,0.00111724,0.0010956466,0.071671836],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005119342,0.00019559005,0.022649983,0.0014939219,0.0000411776,0.000101475714,0.0002008227,0.00031930188,0.00011628258,0.00061643205,0.93053055,0.04322249],"study_design_scores_gemma":[0.0012720615,0.00034861668,0.17182039,0.0025475766,0.000101255755,0.0003573,0.0013474533,0.0014121372,0.0010270235,0.001989601,0.8176251,0.00015148279],"about_ca_topic_score_codex":0.014385053,"about_ca_topic_score_gemma":0.019616354,"teacher_disagreement_score":0.72848916,"about_ca_system_score_codex":0.001359397,"about_ca_system_score_gemma":0.0022416166,"threshold_uncertainty_score":0.90829426},"labels":[],"label_agreement":null},{"id":"W6891760441","doi":"10.48336/sfj9-7954","title":"Comparing information extraction between instance-based data models and relational data models","year":2023,"lang":"en","type":"article","venue":"Memorial University Research Repository (Memorial University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Representation (politics); Complement (music); Information extraction; Data modeling; Data model (GIS); Data extraction; Relational model; External Data Representation","score_opus":0.23689015298217392,"score_gpt":0.3128347531209322,"score_spread":0.07594460013875826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6891760441","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38899454,0.0057693073,0.58039427,0.0033365027,0.00021495434,0.002115063,0.002643521,0.0032537081,0.013278155],"genre_scores_gemma":[0.5302955,0.0023848698,0.4600149,0.0005483982,0.000076675744,0.0009935366,0.0046204273,0.00026951125,0.00079623814],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97883475,0.012519772,0.0025892223,0.0016530378,0.004006787,0.00039636032],"domain_scores_gemma":[0.7663181,0.21210147,0.005187835,0.010695213,0.0051416075,0.0005559257],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027951524,0.0011489087,0.0013096656,0.0051733134,0.00071809476,0.0055366247,0.0018983885,0.0014685637,0.0018344258],"category_scores_gemma":[0.15250935,0.0007195501,0.0022528481,0.0063267844,0.0009493206,0.021816872,0.0036423046,0.001993494,0.00053954596],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049854238,0.0018387511,0.04039115,0.006375973,0.0020518708,0.0006573198,0.009050935,0.040852197,0.010859139,0.06933837,0.008100373,0.8054985],"study_design_scores_gemma":[0.0011874087,0.00373924,0.05113662,0.0025134324,0.0030056958,0.0020605505,0.011306409,0.6732752,0.048368417,0.13924454,0.06351919,0.0006433519],"about_ca_topic_score_codex":0.0022322054,"about_ca_topic_score_gemma":0.0017313279,"teacher_disagreement_score":0.027951524,"about_ca_system_score_codex":0.0019674664,"about_ca_system_score_gemma":0.0013935773,"threshold_uncertainty_score":0.14782357},"labels":[],"label_agreement":null},{"id":"W6891919603","doi":"10.48550/arxiv.cs/0608100","title":"Similarity of Semantic Relations","year":2006,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Analogy; Word (group theory); Similarity (geometry); Vector space model; Relation (database); Vector space; Semantic similarity; Singular value decomposition","score_opus":0.07825484306584243,"score_gpt":0.1825881802614777,"score_spread":0.10433333719563526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6891919603","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3090981,0.01181319,0.5282735,0.0031191893,0.00083123695,0.0011923101,0.014156205,0.0020142102,0.12950204],"genre_scores_gemma":[0.86814106,0.002521442,0.109074116,0.0004881498,0.00041140727,0.00078140583,0.012061619,0.00026436948,0.0062565217],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9920352,0.0020928036,0.000920528,0.0018483305,0.0027727205,0.00033044885],"domain_scores_gemma":[0.99056935,0.0049439925,0.0009982713,0.0016719012,0.0015439866,0.00027244745],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029787987,0.00079738087,0.00089013134,0.012641114,0.0012130926,0.0051404135,0.0012607247,0.0013475745,0.01010365],"category_scores_gemma":[0.028953845,0.00037752578,0.0016419041,0.008964465,0.0024955536,0.00978462,0.003659956,0.0013571022,0.002369157],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005269747,0.00032554235,0.039992273,0.0016276685,0.0009124752,0.00056703197,0.0068609244,0.008023463,0.01214036,0.48219702,0.016270723,0.43055558],"study_design_scores_gemma":[0.00006299753,0.0002198985,0.032083593,0.00034717176,0.00029921113,0.0011819041,0.003836393,0.03535451,0.0034462532,0.84522694,0.07781884,0.00012233523],"about_ca_topic_score_codex":0.0017801244,"about_ca_topic_score_gemma":0.0013114755,"teacher_disagreement_score":0.012641114,"about_ca_system_score_codex":0.0016714344,"about_ca_system_score_gemma":0.0010378913,"threshold_uncertainty_score":0.033800066},"labels":[],"label_agreement":null},{"id":"W6892793348","doi":"10.5281/zenodo.11775378","title":"inférence statistique et probabilités pdf","year":2024,"lang":"fr","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Population; Context (archaeology); Filter (signal processing); Nucleofection","score_opus":0.05388703854162733,"score_gpt":0.27640001846898926,"score_spread":0.22251297992736194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6892793348","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033290356,0.0082291905,0.8675151,0.0146293575,0.0046089543,0.0005736098,0.018559147,0.004571046,0.07798465],"genre_scores_gemma":[0.19087593,0.021421697,0.6032473,0.0073351176,0.011332166,0.0039588204,0.030724924,0.0048836907,0.12622038],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9827068,0.009666513,0.00088188396,0.0027197404,0.0035781453,0.00044689976],"domain_scores_gemma":[0.9052219,0.07770003,0.0023748507,0.009040735,0.005105317,0.00055708626],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.018985974,0.0021457113,0.0032044414,0.006292409,0.0016623283,0.008356621,0.0028234532,0.0031949726,0.13195372],"category_scores_gemma":[0.12896338,0.0020369543,0.0042644725,0.006112503,0.004327126,0.009689296,0.002490701,0.009641831,0.04201851],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020572066,0.00016502538,0.00350092,0.0010384711,0.0005682828,0.00030780004,0.000259582,0.016745208,0.0002604067,0.56640404,0.19744669,0.21309783],"study_design_scores_gemma":[0.00013557383,0.000054880995,0.0019224216,0.0006689959,0.0001274359,0.00031100208,0.00016934275,0.050305016,0.0006810794,0.8277591,0.117769144,0.00009606384],"about_ca_topic_score_codex":0.008567889,"about_ca_topic_score_gemma":0.0058768913,"teacher_disagreement_score":0.8680463,"about_ca_system_score_codex":0.0041638887,"about_ca_system_score_gemma":0.0037566484,"threshold_uncertainty_score":0.44142914},"labels":[],"label_agreement":null},{"id":"W6892862110","doi":"10.5281/zenodo.12803583","title":"IISWC Artifact for \"LLMServingSim: A HW/SW Co-Simulation Infrastructure for LLM Inference Serving at Scale\"","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Artifact (error); Inference; Key (lock); Critical infrastructure; Component (thermodynamics)","score_opus":0.04022454911096261,"score_gpt":0.29112576395574247,"score_spread":0.25090121484477984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6892862110","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054116375,0.00034930013,0.41455686,0.004698559,0.003916614,0.0012696021,0.06689563,0.25550345,0.24739829],"genre_scores_gemma":[0.058310233,0.00042831586,0.2600015,0.0035327873,0.0007171285,0.0019817404,0.24052738,0.17135158,0.26314938],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950852,0.0009105747,0.0001696797,0.00035630044,0.0029571296,0.0005211312],"domain_scores_gemma":[0.9908082,0.0010470774,0.000184601,0.0028355427,0.004065182,0.001059413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056629335,0.0024978644,0.0011240082,0.0013378769,0.0013981415,0.003345528,0.0035155218,0.0021265948,0.13022621],"category_scores_gemma":[0.010963677,0.0012231657,0.0015851913,0.0011924249,0.0008161484,0.002261028,0.0037007648,0.0032125532,0.112125866],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031820944,0.0001107818,0.00048190745,0.00016137205,0.00004751874,0.000066735185,0.00006857756,0.008459962,0.0035556424,0.009777588,0.94353455,0.033417236],"study_design_scores_gemma":[0.0003507262,0.00016484535,0.0007417612,0.000089787914,0.000037578713,0.00009734616,0.00004054715,0.07253156,0.012240538,0.009709689,0.9039192,0.00007645548],"about_ca_topic_score_codex":0.03110735,"about_ca_topic_score_gemma":0.033863008,"teacher_disagreement_score":0.13022621,"about_ca_system_score_codex":0.0026629395,"about_ca_system_score_gemma":0.0055957832,"threshold_uncertainty_score":0.43565005},"labels":[],"label_agreement":null},{"id":"W6893539337","doi":"10.5281/zenodo.16875904","title":"Agentic and Non-Agentic Multi-Hop Systems for Medical Question Answering","year":2025,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Task (project management); Pipeline (software); Question answering; Questions and answers; Joint (building); Semantics (computer science); Interrogative word","score_opus":0.030864254415028207,"score_gpt":0.2675632992729495,"score_spread":0.23669904485792126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6893539337","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026593411,0.0007563605,0.9168474,0.0016738338,0.00019582102,0.0006886756,0.0012517159,0.04601718,0.005975682],"genre_scores_gemma":[0.26878646,0.00025249124,0.71758336,0.0007509224,0.00011941981,0.0005756232,0.003687493,0.0012758733,0.0069684135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99682057,0.0014088497,0.00025652093,0.0007390459,0.0006047189,0.00017024101],"domain_scores_gemma":[0.9935854,0.0035225253,0.00026538293,0.001502982,0.00074434094,0.00037932213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005027283,0.0008877372,0.0008799285,0.0014364271,0.0009886628,0.0029284733,0.0032548,0.002258306,0.008285554],"category_scores_gemma":[0.011575791,0.00063811924,0.0011400572,0.00092162547,0.0009904999,0.0039288313,0.0051847007,0.0019427635,0.003496507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028732365,0.0014310473,0.0066038948,0.001800051,0.00076010113,0.0010860257,0.0026180283,0.121794775,0.050662905,0.066248804,0.07640118,0.6677199],"study_design_scores_gemma":[0.00019004644,0.00015047808,0.00062564696,0.000033273664,0.000064189844,0.00011657738,0.0002376036,0.92544526,0.017025072,0.027325802,0.028730832,0.000055168275],"about_ca_topic_score_codex":0.006256642,"about_ca_topic_score_gemma":0.008570307,"teacher_disagreement_score":0.008285554,"about_ca_system_score_codex":0.0012114714,"about_ca_system_score_gemma":0.002017825,"threshold_uncertainty_score":0.027717948},"labels":[],"label_agreement":null},{"id":"W6894038808","doi":"10.5281/zenodo.8111951","title":"Automated Domain Modeling with Large Language Models: A Comparative Study","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Domain (mathematical analysis); Domain model; Automation; Key (lock); Domain analysis","score_opus":0.07254845640432699,"score_gpt":0.2928745200930867,"score_spread":0.22032606368875968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6894038808","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25424019,0.0031464465,0.69913244,0.0012443619,0.00009703325,0.0007580159,0.0024829654,0.022547524,0.016351031],"genre_scores_gemma":[0.6655793,0.0014427322,0.32329643,0.00016361517,0.000035694924,0.00026099017,0.0055401158,0.0019252946,0.0017558644],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994076,0.0036889145,0.00036598265,0.00048777592,0.0012475767,0.00013379891],"domain_scores_gemma":[0.9536918,0.038038835,0.00091791624,0.005382308,0.0016843774,0.00028477973],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008227366,0.000737645,0.0008618597,0.0023415568,0.00081092806,0.0026649018,0.0020159907,0.0010527375,0.0034151394],"category_scores_gemma":[0.025435999,0.0006073311,0.0013732453,0.0018124188,0.0006793626,0.004778294,0.0021415702,0.0015272014,0.0012046243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018657262,0.0013582524,0.014311714,0.0021968726,0.0006047014,0.00057893933,0.0031465297,0.2198837,0.013090307,0.036286853,0.008930957,0.69774544],"study_design_scores_gemma":[0.00018940485,0.00033927913,0.0043587345,0.00018907693,0.00029551247,0.00038993495,0.0010147607,0.9296418,0.015516791,0.01817601,0.029816683,0.00007204035],"about_ca_topic_score_codex":0.007779112,"about_ca_topic_score_gemma":0.007951584,"teacher_disagreement_score":0.008227366,"about_ca_system_score_codex":0.0016806993,"about_ca_system_score_gemma":0.0017446447,"threshold_uncertainty_score":0.043510973},"labels":[],"label_agreement":null},{"id":"W6901603158","doi":"10.60692/cvwvz-fky71","title":"Guiding the Growth: Difficulty-Controllable Question Generation through Step-by-Step Rewriting","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Rewriting; Interpretability; Task (project management); Inference; Question answering; Control (management)","score_opus":0.050038182334337045,"score_gpt":0.22354264775477928,"score_spread":0.17350446542044223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901603158","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040203933,0.00042715674,0.9506982,0.0006216901,0.00006004326,0.00051275623,0.0004931971,0.0045421408,0.0024408589],"genre_scores_gemma":[0.2747713,0.00018458265,0.7188408,0.0003211482,0.00006219837,0.0003700174,0.0020819674,0.0008216436,0.002546319],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9937848,0.0034581588,0.00031890394,0.0014013132,0.00081127277,0.00022558706],"domain_scores_gemma":[0.97181106,0.022578992,0.0007560878,0.0028605023,0.001635185,0.00035825156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005670236,0.0012969144,0.0008878088,0.0011899633,0.0006206384,0.0016253808,0.0024391958,0.0014015057,0.0041927546],"category_scores_gemma":[0.038093563,0.00057254033,0.0016162398,0.000727851,0.0012142758,0.0032610712,0.0029458392,0.002034667,0.0015426301],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074975775,0.0009192319,0.014291659,0.0017126817,0.00028905986,0.0015017531,0.008051436,0.11639421,0.07965983,0.074081354,0.021235023,0.68111396],"study_design_scores_gemma":[0.00013829432,0.00022421175,0.0015486302,0.00010123099,0.00016161184,0.00045103306,0.0007055849,0.84673357,0.03445257,0.092026845,0.02338586,0.000070526265],"about_ca_topic_score_codex":0.0019423511,"about_ca_topic_score_gemma":0.0034495166,"teacher_disagreement_score":0.005670236,"about_ca_system_score_codex":0.00084830343,"about_ca_system_score_gemma":0.0010749891,"threshold_uncertainty_score":0.029987454},"labels":[],"label_agreement":null},{"id":"W6901609953","doi":"10.60692/fm597-ea287","title":"An Empirical Survey of the Effectiveness of Debiasing Techniques for Pre-trained Language Models","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Debiasing; Counterfactual thinking; Language model; Empirical research; Work (physics); Natural language understanding; Exploit; Cognitive bias","score_opus":0.05684934228093712,"score_gpt":0.29043001269920493,"score_spread":0.2335806704182678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901609953","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60767245,0.056369804,0.29560605,0.0037656573,0.00063841225,0.0007045708,0.0032514,0.020353213,0.011638363],"genre_scores_gemma":[0.77794224,0.0085961465,0.20052996,0.00082944287,0.00022100027,0.00038726174,0.007508436,0.0012523336,0.002733163],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9890616,0.006279328,0.0009644416,0.0017011309,0.0016868361,0.0003065865],"domain_scores_gemma":[0.93804836,0.048375648,0.0018960509,0.0077815633,0.0033673516,0.0005309698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017239925,0.0021583503,0.0011200215,0.0020266725,0.00085008197,0.0014825192,0.002077036,0.0015141915,0.0014837076],"category_scores_gemma":[0.07119617,0.0007053333,0.0013467958,0.0016315806,0.0012463375,0.0039718812,0.0020819134,0.0032849528,0.0013927766],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010221398,0.0004551395,0.0383324,0.0027883383,0.0011874127,0.0001455042,0.0010767815,0.08768324,0.011897857,0.0028124463,0.013831838,0.83876675],"study_design_scores_gemma":[0.00027051035,0.0026382988,0.03934424,0.0016932947,0.0009542181,0.001062648,0.0015691842,0.8260099,0.08001865,0.00962093,0.036552873,0.00026530272],"about_ca_topic_score_codex":0.0053606182,"about_ca_topic_score_gemma":0.007909249,"teacher_disagreement_score":0.017239925,"about_ca_system_score_codex":0.0012655919,"about_ca_system_score_gemma":0.0015222353,"threshold_uncertainty_score":0.09117454},"labels":[],"label_agreement":null},{"id":"W6901614679","doi":"10.60692/7wtz5-ha695","title":"Guiding the Growth: Difficulty-Controllable Question Generation through Step-by-Step Rewriting","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Rewriting; Interpretability; Task (project management); Inference; Question answering; Control (management)","score_opus":0.050038182334337045,"score_gpt":0.22354264775477928,"score_spread":0.17350446542044223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901614679","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040203933,0.00042715674,0.9506982,0.0006216901,0.00006004326,0.00051275623,0.0004931971,0.0045421408,0.0024408589],"genre_scores_gemma":[0.2747713,0.00018458265,0.7188408,0.0003211482,0.00006219837,0.0003700174,0.0020819674,0.0008216436,0.002546319],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937848,0.0034581588,0.00031890394,0.0014013132,0.00081127277,0.00022558706],"domain_scores_gemma":[0.97181106,0.022578992,0.0007560878,0.0028605023,0.001635185,0.00035825156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005670236,0.0012969144,0.0008878088,0.0011899633,0.0006206384,0.0016253808,0.0024391958,0.0014015057,0.0041927546],"category_scores_gemma":[0.038093563,0.00057254033,0.0016162398,0.000727851,0.0012142758,0.0032610712,0.0029458392,0.002034667,0.0015426301],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074975775,0.0009192319,0.014291659,0.0017126817,0.00028905986,0.0015017531,0.008051436,0.11639421,0.07965983,0.074081354,0.021235023,0.68111396],"study_design_scores_gemma":[0.00013829432,0.00022421175,0.0015486302,0.00010123099,0.00016161184,0.00045103306,0.0007055849,0.84673357,0.03445257,0.092026845,0.02338586,0.000070526265],"about_ca_topic_score_codex":0.0019423511,"about_ca_topic_score_gemma":0.0034495166,"teacher_disagreement_score":0.005670236,"about_ca_system_score_codex":0.00084830343,"about_ca_system_score_gemma":0.0010749891,"threshold_uncertainty_score":0.029987454},"labels":[],"label_agreement":null},{"id":"W6901631708","doi":"10.60692/ayq80-vjj26","title":"DeCLUTR: Deep Contrastive Learning for Unsupervised Textual Representations","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; Vector Institute; University of Toronto","funders":"","keywords":"Sentence; Unsupervised learning; Metric (unit); Word (group theory); Cluster analysis; Deep learning; Limiting; Supervised learning","score_opus":0.04129451078165032,"score_gpt":0.23961041336633948,"score_spread":0.19831590258468917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901631708","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012133877,0.00024891517,0.9768077,0.00021151605,0.000080118145,0.00007733996,0.00074408075,0.007917077,0.0017793729],"genre_scores_gemma":[0.32005015,0.0003163687,0.65807784,0.00068683154,0.00010412178,0.00051092857,0.008856658,0.0013205314,0.010076466],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993456,0.00023508651,0.000033500637,0.00022805453,0.00010750376,0.00005032429],"domain_scores_gemma":[0.99885917,0.00054565474,0.000085501975,0.000283801,0.00016658028,0.000059351125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012228348,0.0011711484,0.0005754936,0.00059254083,0.0003883618,0.0008799676,0.0020945498,0.0009977102,0.005108332],"category_scores_gemma":[0.004909346,0.0004932685,0.00074353244,0.000670595,0.0007370836,0.003071238,0.0025866625,0.003052572,0.0030028597],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037895638,0.00036192217,0.002496013,0.000427663,0.00015892298,0.00025745438,0.0004210632,0.16765851,0.029355345,0.044338875,0.037155308,0.7169899],"study_design_scores_gemma":[0.000026564805,0.00013001161,0.00027672425,0.00002462515,0.000013298975,0.00007748666,0.000033854896,0.9562288,0.010154108,0.027330583,0.005684152,0.000019842117],"about_ca_topic_score_codex":0.0019295077,"about_ca_topic_score_gemma":0.005269158,"teacher_disagreement_score":0.005108332,"about_ca_system_score_codex":0.0008830854,"about_ca_system_score_gemma":0.0008081406,"threshold_uncertainty_score":0.017089069},"labels":[],"label_agreement":null},{"id":"W6901662332","doi":"10.60692/9jftz-jtd84","title":"Bridging the Gap between Language Models and Cross-Lingual Sequence Labeling","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Bridging (networking); Language model; Sequence labeling; Margin (machine learning); Regularization (linguistics); Structured prediction; Task (project management); Consistency (knowledge bases); Sequence (biology); Security token","score_opus":0.08578039782984098,"score_gpt":0.2706474233342497,"score_spread":0.1848670255044087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901662332","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06938198,0.0022273352,0.9063878,0.0007892892,0.00026083842,0.00020170453,0.000923997,0.015558154,0.0042689247],"genre_scores_gemma":[0.52794147,0.00076519774,0.4499929,0.0015769161,0.00019269697,0.0006186778,0.008218139,0.0021870574,0.008506957],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99627835,0.0014006875,0.00017524256,0.0016357997,0.0003227258,0.00018719699],"domain_scores_gemma":[0.9927516,0.004076529,0.000324751,0.00183409,0.00079037633,0.00022257505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035659687,0.0023920587,0.0015503583,0.0011940054,0.0009780863,0.0018637219,0.003580983,0.0022522255,0.0034600906],"category_scores_gemma":[0.01163255,0.0011416228,0.0012584248,0.0013504925,0.0014222026,0.003987081,0.0037612338,0.005003884,0.0033706115],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060256454,0.0004189823,0.0041126283,0.0005856162,0.00031474658,0.00037562614,0.0010085126,0.15635978,0.02556601,0.007434434,0.0156056015,0.78761554],"study_design_scores_gemma":[0.000038936698,0.00017600319,0.0014038148,0.000075529075,0.00006939251,0.0001855675,0.0002972889,0.9609789,0.015673393,0.013711118,0.0073373355,0.000052674364],"about_ca_topic_score_codex":0.007320263,"about_ca_topic_score_gemma":0.013128566,"teacher_disagreement_score":0.007320263,"about_ca_system_score_codex":0.0013539484,"about_ca_system_score_gemma":0.0019080671,"threshold_uncertainty_score":0.01885885},"labels":[],"label_agreement":null},{"id":"W6901668656","doi":"10.60692/dv6n0-s9c31","title":"Learning New Skills after Deployment: Improving open-domain internet-driven dialogue with human feedback","year":2023,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Software deployment; Quality (philosophy); Order (exchange); Key (lock); Work (physics); Point (geometry)","score_opus":0.024683373365070253,"score_gpt":0.2190784673812446,"score_spread":0.19439509401617433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901668656","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34109107,0.0019475791,0.6320547,0.0018703537,0.00022599046,0.00037806743,0.00079723855,0.015503787,0.006131212],"genre_scores_gemma":[0.8977744,0.00018688911,0.09628406,0.00053127564,0.000094641444,0.00017780253,0.0014787121,0.0003722576,0.0030999985],"study_design_codex":"simulation_or_modeling","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99774134,0.0013042114,0.00007299393,0.0005237024,0.00017909415,0.00017863119],"domain_scores_gemma":[0.9858504,0.010553325,0.0004687627,0.001583185,0.0009902648,0.0005540507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0055566216,0.0016482695,0.001328484,0.00092834246,0.00061432907,0.0015113609,0.0021323508,0.0018503105,0.0018121117],"category_scores_gemma":[0.022102186,0.00072712306,0.00076191046,0.0006302696,0.00083828054,0.003941148,0.001956934,0.0031164482,0.0017283214],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014419261,0.0018593634,0.015709909,0.00043150556,0.00031044008,0.00017058093,0.0014783096,0.53976077,0.008515023,0.002972371,0.013304924,0.4140448],"study_design_scores_gemma":[0.000043935834,0.00014942793,0.000608143,0.000010891529,0.00002040692,0.000019265317,0.00006372631,0.99529845,0.0015268882,0.0016664334,0.000579024,0.0000134467555],"about_ca_topic_score_codex":0.009884802,"about_ca_topic_score_gemma":0.011467602,"teacher_disagreement_score":0.009884802,"about_ca_system_score_codex":0.0011838524,"about_ca_system_score_gemma":0.0014205923,"threshold_uncertainty_score":0.02938658},"labels":[],"label_agreement":null},{"id":"W6901719222","doi":"10.60692/tr6ct-kqt66","title":"DeCLUTR: Deep Contrastive Learning for Unsupervised Textual Representations","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; Vector Institute; University of Toronto","funders":"","keywords":"Sentence; Unsupervised learning; Metric (unit); Word (group theory); Cluster analysis; Deep learning; Limiting; Supervised learning","score_opus":0.04129451078165032,"score_gpt":0.23961041336633948,"score_spread":0.19831590258468917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901719222","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012133877,0.00024891517,0.9768077,0.00021151605,0.000080118145,0.00007733996,0.00074408075,0.007917077,0.0017793729],"genre_scores_gemma":[0.32005015,0.0003163687,0.65807784,0.00068683154,0.00010412178,0.00051092857,0.008856658,0.0013205314,0.010076466],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993456,0.00023508651,0.000033500637,0.00022805453,0.00010750376,0.00005032429],"domain_scores_gemma":[0.99885917,0.00054565474,0.000085501975,0.000283801,0.00016658028,0.000059351125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012228348,0.0011711484,0.0005754936,0.00059254083,0.0003883618,0.0008799676,0.0020945498,0.0009977102,0.005108332],"category_scores_gemma":[0.004909346,0.0004932685,0.00074353244,0.000670595,0.0007370836,0.003071238,0.0025866625,0.003052572,0.0030028597],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037895638,0.00036192217,0.002496013,0.000427663,0.00015892298,0.00025745438,0.0004210632,0.16765851,0.029355345,0.044338875,0.037155308,0.7169899],"study_design_scores_gemma":[0.000026564805,0.00013001161,0.00027672425,0.00002462515,0.000013298975,0.00007748666,0.000033854896,0.9562288,0.010154108,0.027330583,0.005684152,0.000019842117],"about_ca_topic_score_codex":0.0019295077,"about_ca_topic_score_gemma":0.005269158,"teacher_disagreement_score":0.005108332,"about_ca_system_score_codex":0.0008830854,"about_ca_system_score_gemma":0.0008081406,"threshold_uncertainty_score":0.017089069},"labels":[],"label_agreement":null},{"id":"W6901933573","doi":"10.60692/24n4q-r1960","title":"Proceedings of the Workshop on Generalization in the Age of Deep Learning","year":2018,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; McGill University","funders":"","keywords":"Generalization; Focus (optics); Deep learning; Reading (process); Comprehension; Artificial neural network","score_opus":0.03530438854274401,"score_gpt":0.21786359267141278,"score_spread":0.18255920412866877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901933573","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055329487,0.061168097,0.71816486,0.09135068,0.008398563,0.0002783541,0.0024525325,0.003198123,0.05965933],"genre_scores_gemma":[0.5449563,0.026571613,0.32114303,0.012790849,0.009085091,0.00065433147,0.006366182,0.0015348024,0.076897696],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976745,0.0010305849,0.00015770883,0.0006258222,0.00031752323,0.00019387882],"domain_scores_gemma":[0.98624575,0.009294173,0.00026163794,0.0023280047,0.0012775734,0.00059278327],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010086739,0.0012327965,0.0014040866,0.0008950074,0.0008308174,0.0045455764,0.0028170242,0.0027714309,0.014711802],"category_scores_gemma":[0.024739813,0.0007394525,0.0011928895,0.0009955174,0.0019094225,0.011963093,0.0037476362,0.0059221494,0.0028929696],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007869322,0.00045071266,0.0032597731,0.0006098194,0.00040713913,0.0003328501,0.0010137177,0.041835316,0.0025742834,0.18856812,0.17908895,0.58107233],"study_design_scores_gemma":[0.00010669263,0.0001679896,0.002294179,0.00040304108,0.00014032926,0.00021148665,0.0002837429,0.2516115,0.0032702447,0.5736912,0.16775748,0.0000621609],"about_ca_topic_score_codex":0.0062265364,"about_ca_topic_score_gemma":0.0066532064,"teacher_disagreement_score":0.014711802,"about_ca_system_score_codex":0.0025581599,"about_ca_system_score_gemma":0.001604313,"threshold_uncertainty_score":0.05334443},"labels":[],"label_agreement":null},{"id":"W6902027153","doi":"10.6084/m9.figshare.19976930.v1","title":"Additional file 1 of CoQUAD: a COVID-19 question answering dataset system, facilitating research, benchmarking, and practice","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"","keywords":"Question answering; Exploratory analysis; Exploratory research; Document retrieval","score_opus":0.14672478599036862,"score_gpt":0.37668471526034036,"score_spread":0.22995992926997175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6902027153","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00010726797,0.000008116474,0.00020955273,0.000055054974,0.000012501722,0.000052405478,0.9984622,0.0004104256,0.00068240793],"genre_scores_gemma":[0.0015078811,0.000018942814,0.0018493225,0.00013684986,0.000025198968,0.0008947598,0.99241054,0.00045784944,0.0026986324],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982237,0.0004502155,0.00024794356,0.0005136327,0.00038585256,0.00017850108],"domain_scores_gemma":[0.977749,0.014048974,0.000893379,0.00228942,0.004192611,0.0008266141],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0027132507,0.0010643207,0.0009194648,0.0033592528,0.00095956394,0.0020247195,0.0017712968,0.0013324845,0.59895456],"category_scores_gemma":[0.028870933,0.00051046954,0.0007319312,0.0048354915,0.0003885803,0.0019228377,0.0019299974,0.0012345286,0.22862999],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058796228,0.00003150781,0.00086313253,0.00045395075,0.000009975984,0.000008169883,0.000031688072,0.000085966334,0.00006573935,0.0003396408,0.9950982,0.0029532476],"study_design_scores_gemma":[0.0005640202,0.0000853243,0.0100501655,0.00039087282,0.00003446414,0.000091125694,0.00030232672,0.00090903713,0.00072491285,0.0038176856,0.98296136,0.000068775385],"about_ca_topic_score_codex":0.008031094,"about_ca_topic_score_gemma":0.019888623,"teacher_disagreement_score":0.59895456,"about_ca_system_score_codex":0.0016019737,"about_ca_system_score_gemma":0.0026200837,"threshold_uncertainty_score":0.57204264},"labels":[],"label_agreement":null},{"id":"W6912892384","doi":"10.5281/zenodo.7828789","title":"BudgetLongformer: Can we Cheaply Pretrain a SotA Legal Language Model From Scratch?","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Automatic summarization; Language model; Transformer; Process (computing); Publication; Security token; Task (project management); Code (set theory)","score_opus":0.035854564914102306,"score_gpt":0.2493312381748755,"score_spread":0.21347667326077321,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912892384","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06491794,0.0040631196,0.79769343,0.0054538418,0.0016458223,0.00039832818,0.008109981,0.10098889,0.016728727],"genre_scores_gemma":[0.35449278,0.0018375823,0.5618907,0.00487384,0.00039504864,0.0010109764,0.03504305,0.009038698,0.031417288],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99946326,0.0001504527,0.000036012698,0.00019080064,0.000079873134,0.000079685364],"domain_scores_gemma":[0.99869674,0.0006366511,0.000057295212,0.0003486322,0.00019169606,0.0000690144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016878558,0.0018547968,0.0010557626,0.00073470484,0.00054430973,0.0014393732,0.0029483614,0.001969726,0.01873654],"category_scores_gemma":[0.0069884383,0.0008400181,0.0012565529,0.0007659314,0.000826935,0.005003602,0.0014885071,0.0035967492,0.013041175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010530589,0.0004174052,0.0038703256,0.00082789524,0.00036945398,0.0003508233,0.00023045832,0.23946731,0.01432608,0.015283421,0.15760837,0.56619537],"study_design_scores_gemma":[0.00017431533,0.00016125652,0.00057660794,0.000102591635,0.00008479036,0.00016571124,0.00009548176,0.95001113,0.009307774,0.017404279,0.021866377,0.00004971664],"about_ca_topic_score_codex":0.012630805,"about_ca_topic_score_gemma":0.038581993,"teacher_disagreement_score":0.01873654,"about_ca_system_score_codex":0.0012973895,"about_ca_system_score_gemma":0.0021063907,"threshold_uncertainty_score":0.06267995},"labels":[],"label_agreement":null},{"id":"W6920431187","doi":"10.60692/ys64v-f6230","title":"Balaur: Language Model Pretraining with Lexical Semantic Relations","year":2023,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Generalization; Inference; Set (abstract data type); Meaning (existential); Semantics (computer science); Lexical semantics; Interface (matter); Language model","score_opus":0.041358990059133474,"score_gpt":0.22357447574456035,"score_spread":0.1822154856854269,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920431187","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03291513,0.00071476336,0.9203281,0.00071395264,0.0003291596,0.00018126886,0.0012164867,0.038949866,0.00465132],"genre_scores_gemma":[0.45757905,0.0003802234,0.5176889,0.0012647255,0.000197857,0.0006714777,0.006557065,0.0033269464,0.012333794],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99947745,0.00016785506,0.000028150607,0.0002112635,0.00006345181,0.00005179768],"domain_scores_gemma":[0.9982577,0.0011886106,0.000052608986,0.00026156803,0.00016760577,0.00007194006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013195217,0.0014358672,0.00086520775,0.0006893537,0.00048771207,0.0011831158,0.0023466684,0.0013554881,0.011134171],"category_scores_gemma":[0.005147214,0.0007857011,0.0012405271,0.00059983763,0.0005843591,0.0029145053,0.0021160357,0.00454892,0.0048312233],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000667738,0.0003803659,0.0026468642,0.00044960377,0.0003650934,0.00024546677,0.0005510563,0.18963858,0.021250391,0.01215148,0.04243851,0.72921485],"study_design_scores_gemma":[0.000048640417,0.00006198304,0.00036398336,0.00002608721,0.00003565855,0.000050116512,0.000070746355,0.97876567,0.005468912,0.011092511,0.003992921,0.000022806526],"about_ca_topic_score_codex":0.006038872,"about_ca_topic_score_gemma":0.015268861,"teacher_disagreement_score":0.011134171,"about_ca_system_score_codex":0.00073993637,"about_ca_system_score_gemma":0.0011490849,"threshold_uncertainty_score":0.03724754},"labels":[],"label_agreement":null},{"id":"W6920459592","doi":"10.60692/qkzgq-yt520","title":"Datasets for \"Discourse-Aware Unsupervised Summarization for Long Scientific Documents\"","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Key (lock); Feature (linguistics); Unsupervised learning; Pattern recognition (psychology)","score_opus":0.04695024298544776,"score_gpt":0.2603624005255004,"score_spread":0.21341215754005266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920459592","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057473663,0.0020960267,0.020913191,0.0017516791,0.00043205824,0.0016003521,0.8938043,0.014075802,0.0078529185],"genre_scores_gemma":[0.014795678,0.00027407362,0.03752301,0.00011849582,0.00006373675,0.0013643943,0.94305956,0.0002138595,0.0025872018],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998055,0.0005754615,0.00025847953,0.00044176774,0.00050892273,0.00016039841],"domain_scores_gemma":[0.99438995,0.0019866053,0.00052365963,0.0009256264,0.0017147597,0.00045925728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002187606,0.0016238381,0.00075128814,0.004018606,0.0021131963,0.0011984383,0.0023147396,0.0021861682,0.006311475],"category_scores_gemma":[0.008419218,0.00043320283,0.0014507469,0.0039568823,0.0005770714,0.0013945359,0.0019466672,0.0018876422,0.005468857],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012289747,0.0014995045,0.005531399,0.0054948707,0.0004257243,0.00054416136,0.001642498,0.008452454,0.026841892,0.004516584,0.8138082,0.13001375],"study_design_scores_gemma":[0.0014634421,0.00065276754,0.050102584,0.0006114038,0.00042741868,0.0006972034,0.0023730416,0.043460917,0.045892682,0.0062031667,0.8478026,0.00031270392],"about_ca_topic_score_codex":0.012223218,"about_ca_topic_score_gemma":0.032507185,"teacher_disagreement_score":0.012223218,"about_ca_system_score_codex":0.0015703181,"about_ca_system_score_gemma":0.003069698,"threshold_uncertainty_score":0.024304152},"labels":[],"label_agreement":null},{"id":"W6920460404","doi":"10.60692/hc1k0-3ft41","title":"TIGS: An Inference Algorithm for Text Infilling with Gradient Search","year":2019,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Inference; Missing data; Sequence (biology); Artificial neural network; Generative grammar; Sentence; Generative model; Pattern recognition (psychology)","score_opus":0.04776257973650825,"score_gpt":0.2389968758901704,"score_spread":0.19123429615366216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920460404","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021054274,0.00014321506,0.9925519,0.000089526795,0.000058926653,0.00007191423,0.000099644116,0.004177773,0.0007017499],"genre_scores_gemma":[0.07844155,0.00018725429,0.914832,0.0002642069,0.00013059398,0.00035074406,0.0008842534,0.0014583636,0.0034510726],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915564,0.00024519043,0.00006570367,0.00024929087,0.00021452665,0.00006965138],"domain_scores_gemma":[0.9981712,0.0012243647,0.00009611123,0.00016209115,0.0002757866,0.00007039841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001699032,0.0015228312,0.0013717316,0.0015640424,0.00065718184,0.00133141,0.0028962546,0.0019400009,0.0075237723],"category_scores_gemma":[0.008352464,0.0009372646,0.0011259934,0.0012971368,0.00088818104,0.0024216222,0.0016494321,0.0024486051,0.0039039594],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003344372,0.00015764043,0.0011189175,0.00029529142,0.0001742742,0.00021196055,0.00039912804,0.22711977,0.007935782,0.026817935,0.020825807,0.71460915],"study_design_scores_gemma":[0.00004106671,0.000026381405,0.0000851439,0.000012735043,0.000017259492,0.000038030546,0.0000200104,0.98382485,0.0021175619,0.011134663,0.0026691123,0.000013241415],"about_ca_topic_score_codex":0.0082162265,"about_ca_topic_score_gemma":0.013937251,"teacher_disagreement_score":0.0082162265,"about_ca_system_score_codex":0.00095076807,"about_ca_system_score_gemma":0.0021467693,"threshold_uncertainty_score":0.025169551},"labels":[],"label_agreement":null},{"id":"W6920501266","doi":"10.60692/spg5z-t3n33","title":"Semantics of the Unwritten: The Effect of End of Paragraph and Sequence Tokens on Text Generation with GPT2","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Paragraph; Semantics (computer science); Sequence (biology); Code (set theory); Character (mathematics); Quality (philosophy)","score_opus":0.027984073001474985,"score_gpt":0.20241785085838354,"score_spread":0.17443377785690856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920501266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7688987,0.0023622466,0.18777914,0.00124575,0.0006699676,0.00042252537,0.0024348372,0.02292557,0.0132613685],"genre_scores_gemma":[0.9241624,0.00024159688,0.06518555,0.00036760897,0.000072790775,0.00028866396,0.0043132766,0.0015807974,0.0037873164],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984263,0.0007560492,0.00011111152,0.0004328049,0.00017408808,0.00009951721],"domain_scores_gemma":[0.9895453,0.008173813,0.00033040813,0.00097620685,0.0007208523,0.00025342428],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002766851,0.0015003411,0.00057955866,0.000661392,0.0005174357,0.0012287732,0.0010208816,0.0010331396,0.0036377164],"category_scores_gemma":[0.016385071,0.0004992149,0.00049112324,0.00051089545,0.00054193544,0.0018579733,0.0013451063,0.0020667494,0.0016568232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023444623,0.00054571294,0.023769677,0.001048797,0.00025821768,0.0011517914,0.0013093895,0.20354873,0.09241459,0.0024157036,0.019688256,0.65150464],"study_design_scores_gemma":[0.00031315832,0.00080168445,0.008511089,0.00009407812,0.00020450236,0.00053944846,0.00036223457,0.8924349,0.08458031,0.003931256,0.008152777,0.00007454217],"about_ca_topic_score_codex":0.0036048207,"about_ca_topic_score_gemma":0.0058764904,"teacher_disagreement_score":0.0036377164,"about_ca_system_score_codex":0.00061116496,"about_ca_system_score_gemma":0.0008678867,"threshold_uncertainty_score":0.014632642},"labels":[],"label_agreement":null},{"id":"W6920525879","doi":"10.60692/1h9kf-e4489","title":"QRelScore: Better Evaluating Generated Questions with Deeper Understanding of Context-aware Relevance","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Canadian Institute for Advanced Research","funders":"","keywords":"Relevance (law); Adversarial system; Matching (statistics); Context (archaeology); Metric (unit); Rendering (computer graphics)","score_opus":0.0923810007886917,"score_gpt":0.2493279509683739,"score_spread":0.1569469501796822,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920525879","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2880936,0.0068323235,0.6541147,0.00083590474,0.0006039947,0.0017380695,0.005483459,0.025998915,0.016299004],"genre_scores_gemma":[0.7266391,0.00045560216,0.25982156,0.00039748638,0.00015355054,0.0005592172,0.007556923,0.00079986453,0.003616724],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9920644,0.003794085,0.0005638251,0.0014416118,0.001900898,0.00023517938],"domain_scores_gemma":[0.9729672,0.020248624,0.001516998,0.002320073,0.002304799,0.0006423013],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007915342,0.0018750947,0.0010500501,0.0036919268,0.00049895706,0.0022497494,0.001237513,0.002021905,0.0047504315],"category_scores_gemma":[0.053257767,0.0002376516,0.0007196791,0.0011866089,0.0006821225,0.003306251,0.0022975435,0.0013971746,0.0015212442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021728494,0.00080960407,0.0286116,0.002926109,0.0006054227,0.0003839027,0.0015523841,0.083132915,0.04982031,0.009973132,0.031370748,0.7886411],"study_design_scores_gemma":[0.00026092955,0.0016982985,0.02630575,0.00022232215,0.00026373527,0.0007093057,0.0005152348,0.8778665,0.050414883,0.018730145,0.022753114,0.00025981042],"about_ca_topic_score_codex":0.002213994,"about_ca_topic_score_gemma":0.003573897,"teacher_disagreement_score":0.007915342,"about_ca_system_score_codex":0.00093896996,"about_ca_system_score_gemma":0.0011694696,"threshold_uncertainty_score":0.04186088},"labels":[],"label_agreement":null},{"id":"W6920629036","doi":"10.6084/m9.figshare.26564418","title":"Additional file 1 of Entity and relation extraction from clinical case reports of COVID-19: a natural language processing approach","year":2024,"lang":"en","type":"article","venue":"Open MIND","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"","keywords":"Table (database); Relationship extraction; Natural language; Named-entity recognition; Relation (database); Information extraction; Snippet; Notation","score_opus":0.07103293312264586,"score_gpt":0.38394086096012625,"score_spread":0.3129079278374804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920629036","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00023052956,0.000025120957,0.0006578688,0.00013791646,0.0000357616,0.00014696133,0.99730504,0.0009127026,0.00054804835],"genre_scores_gemma":[0.0045917174,0.00011096299,0.008107842,0.00041631583,0.000104126615,0.0018461901,0.9799589,0.00096970954,0.0038943095],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991341,0.00015995285,0.00022232386,0.00023596657,0.00015745258,0.00009020416],"domain_scores_gemma":[0.98032606,0.014701839,0.0010545278,0.0011101563,0.002294871,0.000512521],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.002304248,0.0016973247,0.0012557582,0.0032372798,0.00079451496,0.0015702544,0.0020578916,0.001472015,0.75380564],"category_scores_gemma":[0.027450502,0.0006057038,0.0011423804,0.0029389411,0.0003229943,0.001865781,0.0017451879,0.0010506916,0.16286834],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044565575,0.00008589052,0.0024294232,0.0031921132,0.000038693994,0.00020767585,0.00008098161,0.0003806686,0.00028959976,0.000397537,0.97826904,0.014182726],"study_design_scores_gemma":[0.0030758118,0.00037653284,0.024929058,0.0037032014,0.00020468202,0.0018170644,0.0009031657,0.0049887444,0.0032953674,0.010615661,0.94587225,0.00021848672],"about_ca_topic_score_codex":0.006355706,"about_ca_topic_score_gemma":0.01332829,"teacher_disagreement_score":0.75380564,"about_ca_system_score_codex":0.001275546,"about_ca_system_score_gemma":0.0025457891,"threshold_uncertainty_score":0.35116637},"labels":[],"label_agreement":null},{"id":"W6924549048","doi":"10.15468/dl.pfpe9r","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Set (abstract data type); Data set","score_opus":0.025341787121003934,"score_gpt":0.2290819036826708,"score_spread":0.20374011656166688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6924549048","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000062958316,0.000034025412,0.000073437375,0.000043437693,0.000014036729,0.000006459553,0.99850273,0.00060505036,0.00065795163],"genre_scores_gemma":[0.000194715,0.00003217165,0.00026198797,0.000040401927,0.0000034546215,0.000042364485,0.9989405,0.0001295902,0.00035488265],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99913055,0.00012663576,0.0001083578,0.00030274165,0.00019388532,0.00013783357],"domain_scores_gemma":[0.9980964,0.0005550656,0.00016976375,0.0005126118,0.00043021826,0.00023593381],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00090045866,0.0020600548,0.0013643847,0.0041293376,0.00093917036,0.0023067344,0.002817268,0.0020877179,0.09873036],"category_scores_gemma":[0.005256482,0.0008041564,0.0012906792,0.008165091,0.00039973197,0.0021332735,0.0021809787,0.0020377473,0.16251382],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028257047,0.000013910986,0.00043040753,0.00047173654,0.000015732816,0.000016721546,0.000025068617,0.00019099693,0.00012584314,0.0004289465,0.9966252,0.0016272902],"study_design_scores_gemma":[0.00007744422,0.000008937767,0.0019515123,0.00016222506,0.000015485195,0.000045137287,0.00007840756,0.00033262873,0.00024388212,0.0009922617,0.99607325,0.000018789839],"about_ca_topic_score_codex":0.021704515,"about_ca_topic_score_gemma":0.038532093,"teacher_disagreement_score":0.9012696,"about_ca_system_score_codex":0.001668992,"about_ca_system_score_gemma":0.0022577269,"threshold_uncertainty_score":0.33028597},"labels":[],"label_agreement":null},{"id":"W6926576609","doi":"10.25316/ir-13083","title":"The Nanaimo Free Press [Tuesday, October 29, 1889]","year":2019,"lang":"en","type":"other","venue":"VIURRSpace (Vancouver Island University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.010664930351368127,"score_gpt":0.19592328766186343,"score_spread":0.1852583573104953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6926576609","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008238359,0.032251123,0.0006059919,0.007967803,0.0073126187,0.000029011966,0.00407392,0.00037094255,0.94656473],"genre_scores_gemma":[0.0020741306,0.0033328813,0.00013845305,0.00018916096,0.00035601537,0.000010572367,0.0004370674,0.00012462777,0.99333704],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99952316,0.00005685288,0.00002022934,0.00009704038,0.00023588675,0.00006675781],"domain_scores_gemma":[0.99957365,0.000087168715,0.000028310847,0.00004206393,0.00018071874,0.000088042114],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005366689,0.0008946881,0.00069311616,0.0024685406,0.0033656864,0.007379802,0.0007204802,0.0015795775,0.23582213],"category_scores_gemma":[0.0024173178,0.0004255664,0.00033502548,0.0045753103,0.0008921445,0.0025640102,0.0013644997,0.0018202272,0.083717436],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000023011822,0.0000060212715,0.00016102061,0.0000845764,0.0000042982497,0.000053060638,0.00013112347,0.00008554367,0.00005926318,0.017204523,0.94350916,0.038678315],"study_design_scores_gemma":[0.0000013239945,0.000001589906,0.00030584986,0.00005701891,9.64231e-7,0.000015976495,0.000041166903,0.000027161746,0.000025148986,0.00067871483,0.9988426,0.0000023933965],"about_ca_topic_score_codex":0.097705625,"about_ca_topic_score_gemma":0.35704592,"teacher_disagreement_score":0.9022944,"about_ca_system_score_codex":0.0054191197,"about_ca_system_score_gemma":0.002957114,"threshold_uncertainty_score":0.7889036},"labels":[],"label_agreement":null},{"id":"W6931653815","doi":"10.5281/zenodo.7828790","title":"BudgetLongformer: Can we Cheaply Pretrain a SotA Legal Language Model From Scratch?","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Automatic summarization; Language model; Transformer; Process (computing); Publication; Security token; Task (project management); Code (set theory)","score_opus":0.035854564914102306,"score_gpt":0.2493312381748755,"score_spread":0.21347667326077321,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931653815","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06491794,0.0040631196,0.79769343,0.0054538418,0.0016458223,0.00039832818,0.008109981,0.10098889,0.016728727],"genre_scores_gemma":[0.35449278,0.0018375823,0.5618907,0.00487384,0.00039504864,0.0010109764,0.03504305,0.009038698,0.031417288],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99946326,0.0001504527,0.000036012698,0.00019080064,0.000079873134,0.000079685364],"domain_scores_gemma":[0.99869674,0.0006366511,0.000057295212,0.0003486322,0.00019169606,0.0000690144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016878558,0.0018547968,0.0010557626,0.00073470484,0.00054430973,0.0014393732,0.0029483614,0.001969726,0.01873654],"category_scores_gemma":[0.0069884383,0.0008400181,0.0012565529,0.0007659314,0.000826935,0.005003602,0.0014885071,0.0035967492,0.013041175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010530589,0.0004174052,0.0038703256,0.00082789524,0.00036945398,0.0003508233,0.00023045832,0.23946731,0.01432608,0.015283421,0.15760837,0.56619537],"study_design_scores_gemma":[0.00017431533,0.00016125652,0.00057660794,0.000102591635,0.00008479036,0.00016571124,0.00009548176,0.95001113,0.009307774,0.017404279,0.021866377,0.00004971664],"about_ca_topic_score_codex":0.012630805,"about_ca_topic_score_gemma":0.038581993,"teacher_disagreement_score":0.01873654,"about_ca_system_score_codex":0.0012973895,"about_ca_system_score_gemma":0.0021063907,"threshold_uncertainty_score":0.06267995},"labels":[],"label_agreement":null},{"id":"W6931697980","doi":"10.5281/zenodo.7497955","title":"TuringLang/Turing.jl: v0.23.3","year":2023,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Nucleofection; Gestational period; TSG101; Diafiltration; Dysgeusia; Liquation; Fusible alloy; Demotion; Emperipolesis","score_opus":0.042686971195317176,"score_gpt":0.2429998452374858,"score_spread":0.20031287404216863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931697980","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00042481575,0.0006913765,0.07720407,0.0024114847,0.0009319393,0.00017297606,0.030867638,0.5280748,0.3592209],"genre_scores_gemma":[0.020704038,0.0013637147,0.036083277,0.0026736765,0.0005338014,0.0007971842,0.04977599,0.5430612,0.34500718],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981529,0.000341727,0.00012396518,0.00037218526,0.0006721112,0.00033703304],"domain_scores_gemma":[0.9969693,0.0005644935,0.000101308055,0.0012789469,0.00078181614,0.0003042407],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0019427903,0.0027348825,0.0019633556,0.0019777215,0.001657878,0.0074579427,0.0053778426,0.0035379587,0.6416749],"category_scores_gemma":[0.010079957,0.0031085385,0.0021887294,0.002428272,0.0015470829,0.008729669,0.006405089,0.0058428957,0.7747435],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006677066,0.000015130441,0.00013139758,0.00027305022,0.000011214421,0.00003758494,0.00007011179,0.00028865307,0.0003825062,0.020849414,0.9509605,0.026913747],"study_design_scores_gemma":[0.00003136367,0.000009890301,0.0002020104,0.00012663026,0.000007848182,0.00009617128,0.000031901604,0.0011284016,0.0018801879,0.015632397,0.98080736,0.00004581842],"about_ca_topic_score_codex":0.005621241,"about_ca_topic_score_gemma":0.0039301584,"teacher_disagreement_score":0.6416749,"about_ca_system_score_codex":0.0032338845,"about_ca_system_score_gemma":0.0020991785,"threshold_uncertainty_score":0.51110727},"labels":[],"label_agreement":null},{"id":"W6939191839","doi":"10.60692/rwgra-g5d61","title":"The Power of Prompt Tuning for Low-Resource Semantic Parsing","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Task (project management); Natural language; Meaning (existential); Parsing; Language model; Natural language understanding; Semantics (computer science); Natural language generation","score_opus":0.029007153236706978,"score_gpt":0.20894531611168146,"score_spread":0.1799381628749745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939191839","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23922078,0.0026251592,0.64056474,0.0024433658,0.00086781156,0.00037962027,0.002414299,0.09774375,0.013740376],"genre_scores_gemma":[0.82685095,0.00039473976,0.16012546,0.001246136,0.00014120564,0.00031739354,0.003848427,0.0040564,0.003019222],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99791604,0.00087112683,0.000121696015,0.0006896108,0.00023365446,0.00016786637],"domain_scores_gemma":[0.9933276,0.0044575515,0.00017801372,0.0013420169,0.00045902247,0.00023581502],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042424933,0.0020196768,0.0009370686,0.000704263,0.0007610056,0.0020935286,0.002173456,0.001875249,0.0054912507],"category_scores_gemma":[0.024897566,0.0006078398,0.0009020908,0.000912695,0.0009874734,0.0050996486,0.0023119266,0.0047634775,0.0027311137],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002383856,0.00095060735,0.013729865,0.0009050155,0.00034386502,0.0006998562,0.0010742734,0.34462237,0.053379476,0.01547678,0.05050574,0.51592827],"study_design_scores_gemma":[0.0002849255,0.00042741594,0.0021444142,0.00006499777,0.000088108616,0.00021763917,0.00036011913,0.9353479,0.019592479,0.029746529,0.011649097,0.00007643702],"about_ca_topic_score_codex":0.0050209165,"about_ca_topic_score_gemma":0.0077414257,"teacher_disagreement_score":0.0054912507,"about_ca_system_score_codex":0.0010367635,"about_ca_system_score_gemma":0.0022345856,"threshold_uncertainty_score":0.022436738},"labels":[],"label_agreement":null},{"id":"W6939224053","doi":"10.60692/w9cgp-8cn47","title":"Feeding What You Need by Understanding What You Learned","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Pipeline (software); Interpretation (philosophy); Comprehension; Training set; Reading (process); Curriculum; Deep learning","score_opus":0.08032085247982688,"score_gpt":0.22488749781583378,"score_spread":0.14456664533600688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939224053","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23133262,0.0011440978,0.72229993,0.005410756,0.00031911177,0.00060220214,0.004202355,0.01239589,0.022292992],"genre_scores_gemma":[0.5394021,0.0011280014,0.43843746,0.00087740034,0.00012177066,0.0006084441,0.007920262,0.00070227374,0.010802225],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990031,0.00037138988,0.000056486366,0.00036083127,0.00014982959,0.00005843822],"domain_scores_gemma":[0.99446243,0.0037743873,0.0002574386,0.000761646,0.00054599636,0.00019797451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015534571,0.001514397,0.00081097323,0.001229219,0.0005218633,0.0024311198,0.0013875914,0.0018384807,0.013536249],"category_scores_gemma":[0.018382104,0.00059620856,0.001135781,0.0010862863,0.0004768337,0.0068519562,0.002497848,0.002620915,0.005926637],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004591481,0.0010010756,0.024561157,0.0014882841,0.00017702261,0.00043166676,0.0038107017,0.03701868,0.030689504,0.014761017,0.032484826,0.85311687],"study_design_scores_gemma":[0.00011616268,0.00085884763,0.016874518,0.00058307115,0.00026302965,0.0006170107,0.0035467779,0.7312581,0.042135485,0.121650524,0.08188514,0.00021131034],"about_ca_topic_score_codex":0.0029669912,"about_ca_topic_score_gemma":0.0045234263,"teacher_disagreement_score":0.013536249,"about_ca_system_score_codex":0.0006579085,"about_ca_system_score_gemma":0.0014468686,"threshold_uncertainty_score":0.045283258},"labels":[],"label_agreement":null},{"id":"W6939282806","doi":"10.60692/7v46n-wbe56","title":"Datasets for \"Discourse-Aware Unsupervised Summarization for Long Scientific Documents\"","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Key (lock); Feature (linguistics); Unsupervised learning; Pattern recognition (psychology)","score_opus":0.04695024298544776,"score_gpt":0.2603624005255004,"score_spread":0.21341215754005266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939282806","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057473663,0.0020960267,0.020913191,0.0017516791,0.00043205824,0.0016003521,0.8938043,0.014075802,0.0078529185],"genre_scores_gemma":[0.014795678,0.00027407362,0.03752301,0.00011849582,0.00006373675,0.0013643943,0.94305956,0.0002138595,0.0025872018],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998055,0.0005754615,0.00025847953,0.00044176774,0.00050892273,0.00016039841],"domain_scores_gemma":[0.99438995,0.0019866053,0.00052365963,0.0009256264,0.0017147597,0.00045925728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002187606,0.0016238381,0.00075128814,0.004018606,0.0021131963,0.0011984383,0.0023147396,0.0021861682,0.006311475],"category_scores_gemma":[0.008419218,0.00043320283,0.0014507469,0.0039568823,0.0005770714,0.0013945359,0.0019466672,0.0018876422,0.005468857],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012289747,0.0014995045,0.005531399,0.0054948707,0.0004257243,0.00054416136,0.001642498,0.008452454,0.026841892,0.004516584,0.8138082,0.13001375],"study_design_scores_gemma":[0.0014634421,0.00065276754,0.050102584,0.0006114038,0.00042741868,0.0006972034,0.0023730416,0.043460917,0.045892682,0.0062031667,0.8478026,0.00031270392],"about_ca_topic_score_codex":0.012223218,"about_ca_topic_score_gemma":0.032507185,"teacher_disagreement_score":0.012223218,"about_ca_system_score_codex":0.0015703181,"about_ca_system_score_gemma":0.003069698,"threshold_uncertainty_score":0.024304152},"labels":[],"label_agreement":null},{"id":"W6939456355","doi":"10.60692/xvrwf-nqj55","title":"Datasets for \"Discourse-Aware Unsupervised Summarization for Long Scientific Documents\"","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Key (lock); Feature (linguistics); Unsupervised learning; Pattern recognition (psychology)","score_opus":0.04695024298544776,"score_gpt":0.2603624005255004,"score_spread":0.21341215754005266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939456355","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057473663,0.0020960267,0.020913191,0.0017516791,0.00043205824,0.0016003521,0.8938043,0.014075802,0.0078529185],"genre_scores_gemma":[0.014795678,0.00027407362,0.03752301,0.00011849582,0.00006373675,0.0013643943,0.94305956,0.0002138595,0.0025872018],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998055,0.0005754615,0.00025847953,0.00044176774,0.00050892273,0.00016039841],"domain_scores_gemma":[0.99438995,0.0019866053,0.00052365963,0.0009256264,0.0017147597,0.00045925728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002187606,0.0016238381,0.00075128814,0.004018606,0.0021131963,0.0011984383,0.0023147396,0.0021861682,0.006311475],"category_scores_gemma":[0.008419218,0.00043320283,0.0014507469,0.0039568823,0.0005770714,0.0013945359,0.0019466672,0.0018876422,0.005468857],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012289747,0.0014995045,0.005531399,0.0054948707,0.0004257243,0.00054416136,0.001642498,0.008452454,0.026841892,0.004516584,0.8138082,0.13001375],"study_design_scores_gemma":[0.0014634421,0.00065276754,0.050102584,0.0006114038,0.00042741868,0.0006972034,0.0023730416,0.043460917,0.045892682,0.0062031667,0.8478026,0.00031270392],"about_ca_topic_score_codex":0.012223218,"about_ca_topic_score_gemma":0.032507185,"teacher_disagreement_score":0.012223218,"about_ca_system_score_codex":0.0015703181,"about_ca_system_score_gemma":0.003069698,"threshold_uncertainty_score":0.024304152},"labels":[],"label_agreement":null},{"id":"W6949396090","doi":"10.5281/zenodo.14167444","title":"Pronoun Generation for Text Summarization and Question Answering","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Question answering; Automatic summarization; Pronoun; Natural language; Corpus linguistics","score_opus":0.016920509189539544,"score_gpt":0.23479597512023376,"score_spread":0.2178754659306942,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949396090","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012454803,0.0032901412,0.94351923,0.0014378532,0.000825544,0.00059668644,0.0048311125,0.025428617,0.007616017],"genre_scores_gemma":[0.15000196,0.0016262397,0.81363976,0.00034692005,0.00075993687,0.00059871847,0.01823536,0.0016465319,0.013144584],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976457,0.001190764,0.00017203936,0.00047880044,0.0003372229,0.00017531613],"domain_scores_gemma":[0.99609154,0.0020671217,0.00020332047,0.00059508125,0.00092395354,0.00011902908],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026218088,0.0013295183,0.0013522344,0.0035219307,0.0013634936,0.0020223304,0.0016238574,0.0015087373,0.01784298],"category_scores_gemma":[0.008728143,0.0005992656,0.0009924343,0.0029617178,0.00052283896,0.0032432403,0.001721613,0.0012124353,0.0133886915],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003839078,0.00016271157,0.0007370857,0.0009841344,0.00011657652,0.00025471445,0.00060790876,0.008329427,0.027683744,0.016856654,0.08603294,0.8578504],"study_design_scores_gemma":[0.0002036446,0.00037171718,0.0018412273,0.00019376645,0.00024798876,0.0005769933,0.0010202944,0.5834906,0.08494214,0.13082403,0.19617492,0.00011264742],"about_ca_topic_score_codex":0.0018706421,"about_ca_topic_score_gemma":0.0030257902,"teacher_disagreement_score":0.01784298,"about_ca_system_score_codex":0.00067365036,"about_ca_system_score_gemma":0.001022304,"threshold_uncertainty_score":0.059690714},"labels":[],"label_agreement":null},{"id":"W6949444982","doi":"10.5281/zenodo.14167443","title":"Pronoun Generation for Text Summarization and Question Answering","year":2006,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Question answering; Automatic summarization; Pronoun; Natural language; Corpus linguistics","score_opus":0.03167901767740509,"score_gpt":0.23481886717501124,"score_spread":0.20313984949760616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949444982","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0111557115,0.0031336744,0.94478923,0.0014129037,0.00068915397,0.0006396675,0.0051791132,0.026385507,0.006614995],"genre_scores_gemma":[0.13592821,0.0015441788,0.8291373,0.00033685967,0.00067957834,0.0006999773,0.01914825,0.0016585456,0.01086705],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970892,0.0015189996,0.00021556624,0.0005970095,0.00038151606,0.00019761821],"domain_scores_gemma":[0.99509984,0.0027895248,0.00025344276,0.0007068035,0.0010165996,0.00013374422],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031759162,0.0015353542,0.0014808071,0.0040486543,0.0014502903,0.0021978824,0.001753035,0.0017339761,0.0178897],"category_scores_gemma":[0.010038408,0.0006764444,0.0011339794,0.0033297832,0.00058025843,0.003635311,0.0019556633,0.0013267442,0.014201103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041325146,0.00016310676,0.0007474145,0.0011794735,0.00013630498,0.0002540797,0.0007363095,0.008999197,0.026102224,0.017677005,0.08312779,0.8604638],"study_design_scores_gemma":[0.00020677007,0.00034924477,0.0016527943,0.0002203629,0.00025753697,0.00054680125,0.0011089819,0.5960758,0.07208115,0.14346959,0.18391724,0.00011366364],"about_ca_topic_score_codex":0.0019892391,"about_ca_topic_score_gemma":0.0030034475,"teacher_disagreement_score":0.0178897,"about_ca_system_score_codex":0.0007838659,"about_ca_system_score_gemma":0.0011298653,"threshold_uncertainty_score":0.059847057},"labels":[],"label_agreement":null},{"id":"W6950165384","doi":"10.5281/zenodo.6798019","title":"Touché21-Argument-Retrieval-for-Comparative-Questions","year":2024,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Optech (Canada)","funders":"","keywords":"Task (project management); Argument (complex analysis); Relation (database); Matching (statistics); Product (mathematics)","score_opus":0.06724454125565449,"score_gpt":0.29595275756642736,"score_spread":0.22870821631077287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6950165384","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054396175,0.0015086032,0.018314946,0.0020183849,0.0007227842,0.0017864613,0.8542229,0.020474084,0.046555635],"genre_scores_gemma":[0.059924394,0.00016272523,0.015815573,0.00046249243,0.00014193906,0.003322529,0.9046718,0.002646069,0.012852454],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951383,0.0017018126,0.00050686015,0.000902876,0.0013757647,0.00037433326],"domain_scores_gemma":[0.9706143,0.019652361,0.00086753786,0.0052606864,0.00252167,0.0010834176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033498837,0.0027532736,0.0018282277,0.0031368271,0.0018734863,0.0025914626,0.0021968014,0.005533522,0.10693024],"category_scores_gemma":[0.03600607,0.0009620279,0.0016334956,0.0018372796,0.00087030814,0.00520646,0.005507955,0.0037194886,0.099822514],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024602306,0.0010274362,0.005143059,0.0024030174,0.0001571683,0.00050705473,0.0010810804,0.0017188368,0.0069651045,0.0048458865,0.92809993,0.04559125],"study_design_scores_gemma":[0.0030118956,0.00077109045,0.033137698,0.0004492436,0.0001081807,0.0013645282,0.0015356903,0.016276393,0.021089202,0.017184958,0.9047721,0.00029895324],"about_ca_topic_score_codex":0.005711526,"about_ca_topic_score_gemma":0.008607161,"teacher_disagreement_score":0.10693024,"about_ca_system_score_codex":0.0011896142,"about_ca_system_score_gemma":0.0014259402,"threshold_uncertainty_score":0.35771728},"labels":[],"label_agreement":null},{"id":"W6950639761","doi":"10.5683/sp3/a7itkl","title":"Replication Data for: Comparison of four competing invasion percolation models for gas flow in porous media","year":2023,"lang":"en","type":"dataset","venue":"Borealis","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Percolation (cognitive psychology); Porous medium; Replication (statistics); Flow (mathematics); Vadose zone; Percolation theory; Porosity","score_opus":0.19114104330353726,"score_gpt":0.35323096001273807,"score_spread":0.1620899167092008,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6950639761","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001523483,0.00031296813,0.00045091528,0.00019398915,0.00009610818,0.000063885724,0.994951,0.0011834798,0.0012240976],"genre_scores_gemma":[0.0035520967,0.00010110507,0.0012717308,0.00007718835,0.000025248417,0.00025464225,0.99373984,0.000117397394,0.00086077105],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985203,0.00030050555,0.00016352997,0.00046469408,0.00037074907,0.00018015785],"domain_scores_gemma":[0.9958988,0.0016911654,0.00026867323,0.0010783754,0.0008468052,0.0002161895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020612981,0.0026725447,0.0015403009,0.0027204314,0.0011256696,0.002343905,0.0035982495,0.0030614152,0.034498],"category_scores_gemma":[0.01038725,0.0005340228,0.002391136,0.0027952616,0.0006391375,0.0011807791,0.0018896249,0.0019200268,0.050118934],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000254553,0.00015117392,0.0032265035,0.0017134883,0.00011999088,0.0000778118,0.00007212624,0.0015926278,0.0005967202,0.0007085099,0.98210573,0.00938076],"study_design_scores_gemma":[0.001153536,0.0001533262,0.024823502,0.00072020927,0.00022272862,0.0004786112,0.00040801466,0.009061642,0.0029828965,0.0055269175,0.95431674,0.00015178809],"about_ca_topic_score_codex":0.016855404,"about_ca_topic_score_gemma":0.03670552,"teacher_disagreement_score":0.034498,"about_ca_system_score_codex":0.0015321091,"about_ca_system_score_gemma":0.0017444587,"threshold_uncertainty_score":0.11540735},"labels":[],"label_agreement":null},{"id":"W6957386035","doi":"10.60692/kk5ze-cry69","title":"Semantics of the Unwritten: The Effect of End of Paragraph and Sequence Tokens on Text Generation with GPT2","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Paragraph; Semantics (computer science); Sequence (biology); Code (set theory); Character (mathematics); Quality (philosophy)","score_opus":0.027984073001474985,"score_gpt":0.20241785085838354,"score_spread":0.17443377785690856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6957386035","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7688987,0.0023622466,0.18777914,0.00124575,0.0006699676,0.00042252537,0.0024348372,0.02292557,0.0132613685],"genre_scores_gemma":[0.9241624,0.00024159688,0.06518555,0.00036760897,0.000072790775,0.00028866396,0.0043132766,0.0015807974,0.0037873164],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984263,0.0007560492,0.00011111152,0.0004328049,0.00017408808,0.00009951721],"domain_scores_gemma":[0.9895453,0.008173813,0.00033040813,0.00097620685,0.0007208523,0.00025342428],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002766851,0.0015003411,0.00057955866,0.000661392,0.0005174357,0.0012287732,0.0010208816,0.0010331396,0.0036377164],"category_scores_gemma":[0.016385071,0.0004992149,0.00049112324,0.00051089545,0.00054193544,0.0018579733,0.0013451063,0.0020667494,0.0016568232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023444623,0.00054571294,0.023769677,0.001048797,0.00025821768,0.0011517914,0.0013093895,0.20354873,0.09241459,0.0024157036,0.019688256,0.65150464],"study_design_scores_gemma":[0.00031315832,0.00080168445,0.008511089,0.00009407812,0.00020450236,0.00053944846,0.00036223457,0.8924349,0.08458031,0.003931256,0.008152777,0.00007454217],"about_ca_topic_score_codex":0.0036048207,"about_ca_topic_score_gemma":0.0058764904,"teacher_disagreement_score":0.0036377164,"about_ca_system_score_codex":0.00061116496,"about_ca_system_score_gemma":0.0008678867,"threshold_uncertainty_score":0.014632642},"labels":[],"label_agreement":null},{"id":"W6957604568","doi":"10.60692/adptg-twe02","title":"On-the-Fly Attention Modulation for Neural Generation","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Modulation (music); Artificial neural network; Neural activity; Prior probability; Computation; Deep neural networks; Simple (philosophy)","score_opus":0.07389547561211676,"score_gpt":0.2246698726626971,"score_spread":0.15077439705058035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6957604568","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13565348,0.00137626,0.8435111,0.00060142466,0.00020301345,0.00014432485,0.00030580373,0.010603523,0.0076011214],"genre_scores_gemma":[0.8767486,0.00024410513,0.11946398,0.00021237224,0.00006821641,0.00012555439,0.00024672944,0.00029389653,0.00259649],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.999762,0.000082226245,0.000014879262,0.00006685494,0.000046680198,0.000027276294],"domain_scores_gemma":[0.999303,0.00042857527,0.00005254633,0.00011962153,0.000064256776,0.00003204416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00046341636,0.0005631474,0.00031649633,0.00029104634,0.00019741303,0.00047146843,0.0010310615,0.00059145154,0.0034576517],"category_scores_gemma":[0.0027442428,0.00020160127,0.0002680534,0.00026616023,0.0002987478,0.0010422688,0.0008277707,0.00097604486,0.00065711985],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003131087,0.0002213983,0.0013631383,0.00020129533,0.000067003966,0.0001637564,0.00021832116,0.10136976,0.08451644,0.008542587,0.0048594647,0.79816353],"study_design_scores_gemma":[0.00003705117,0.00008516411,0.00069081323,0.000012402897,0.000022481972,0.00004641287,0.000023781186,0.95891595,0.028739328,0.009107614,0.002308728,0.000010280303],"about_ca_topic_score_codex":0.0018407282,"about_ca_topic_score_gemma":0.0034657957,"teacher_disagreement_score":0.0034576517,"about_ca_system_score_codex":0.0005566823,"about_ca_system_score_gemma":0.00047147044,"threshold_uncertainty_score":0.011567056},"labels":[],"label_agreement":null},{"id":"W6957621777","doi":"10.60692/5r3jg-m5e18","title":"Imperfect also Deserves Reward: Multi-Level and Sequential Reward Modeling for Better Dialog Management","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Imperfect; Dialog box; Selection (genetic algorithm); Language model; Association (psychology); Subject (documents)","score_opus":0.09650262403361097,"score_gpt":0.25009867369707384,"score_spread":0.15359604966346285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6957621777","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06339035,0.0014746774,0.9225486,0.0031014963,0.00036942816,0.000147796,0.000787402,0.0015345982,0.006645704],"genre_scores_gemma":[0.91925496,0.0003970747,0.07324536,0.0003117552,0.00020159165,0.00019365795,0.0004916941,0.00020709482,0.0056968736],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859077,0.0006441832,0.00006283429,0.00032750334,0.00014781636,0.00022676168],"domain_scores_gemma":[0.9942577,0.004044933,0.00035766704,0.00034688984,0.0004937782,0.00049898215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039448654,0.0013282171,0.0019007997,0.0009347497,0.0012654484,0.0025039527,0.0025640079,0.0019949623,0.008843505],"category_scores_gemma":[0.013475849,0.0008568039,0.0012522861,0.00071726274,0.0009571686,0.0042702877,0.0025238,0.0034231532,0.0011595186],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010806021,0.00035622105,0.0065499274,0.00021310769,0.00022024392,0.0002684245,0.00083393906,0.76648504,0.00170433,0.098561525,0.009925647,0.113800935],"study_design_scores_gemma":[0.000014657571,0.00002119901,0.00019441669,0.000007572637,0.00001856396,0.000010314836,0.00001636078,0.9811126,0.00009519466,0.018044798,0.00045563024,0.000008678594],"about_ca_topic_score_codex":0.015111044,"about_ca_topic_score_gemma":0.022119626,"teacher_disagreement_score":0.015111044,"about_ca_system_score_codex":0.0019111417,"about_ca_system_score_gemma":0.0027085494,"threshold_uncertainty_score":0.030046165},"labels":[],"label_agreement":null},{"id":"W6957827635","doi":"10.60692/qzwgz-q7g37","title":"TIGS: An Inference Algorithm for Text Infilling with Gradient Search","year":2019,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Inference; Missing data; Sequence (biology); Artificial neural network; Generative grammar; Sentence; Generative model; Pattern recognition (psychology)","score_opus":0.04776257973650825,"score_gpt":0.2389968758901704,"score_spread":0.19123429615366216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6957827635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021054274,0.00014321506,0.9925519,0.000089526795,0.000058926653,0.00007191423,0.000099644116,0.004177773,0.0007017499],"genre_scores_gemma":[0.07844155,0.00018725429,0.914832,0.0002642069,0.00013059398,0.00035074406,0.0008842534,0.0014583636,0.0034510726],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915564,0.00024519043,0.00006570367,0.00024929087,0.00021452665,0.00006965138],"domain_scores_gemma":[0.9981712,0.0012243647,0.00009611123,0.00016209115,0.0002757866,0.00007039841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001699032,0.0015228312,0.0013717316,0.0015640424,0.00065718184,0.00133141,0.0028962546,0.0019400009,0.0075237723],"category_scores_gemma":[0.008352464,0.0009372646,0.0011259934,0.0012971368,0.00088818104,0.0024216222,0.0016494321,0.0024486051,0.0039039594],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003344372,0.00015764043,0.0011189175,0.00029529142,0.0001742742,0.00021196055,0.00039912804,0.22711977,0.007935782,0.026817935,0.020825807,0.71460915],"study_design_scores_gemma":[0.00004106671,0.000026381405,0.0000851439,0.000012735043,0.000017259492,0.000038030546,0.0000200104,0.98382485,0.0021175619,0.011134663,0.0026691123,0.000013241415],"about_ca_topic_score_codex":0.0082162265,"about_ca_topic_score_gemma":0.013937251,"teacher_disagreement_score":0.0082162265,"about_ca_system_score_codex":0.00095076807,"about_ca_system_score_gemma":0.0021467693,"threshold_uncertainty_score":0.025169551},"labels":[],"label_agreement":null},{"id":"W6957997596","doi":"10.60692/23tr5-rj647","title":"Learning to Transfer Prompts for Text Generation","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Text generation; Set (abstract data type); Reuse; Transfer of learning; Training set; Mechanism (biology); Labeled data; Language model","score_opus":0.05304766468987507,"score_gpt":0.21871948864084198,"score_spread":0.1656718239509669,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6957997596","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038965434,0.0007197615,0.9417918,0.00040453856,0.00032503353,0.0002939364,0.00050748605,0.0146733485,0.0023187688],"genre_scores_gemma":[0.5932238,0.0004579387,0.39518747,0.00061566726,0.00025733063,0.0008222867,0.0021778843,0.0009618146,0.0062957513],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987937,0.0004799575,0.0000603851,0.00047971634,0.00012533508,0.000060886574],"domain_scores_gemma":[0.99560475,0.0028990759,0.00023750823,0.0006393633,0.0004261648,0.00019318693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023040103,0.0015298292,0.0007648949,0.00064569386,0.00034570013,0.0008112932,0.0016862634,0.0013864592,0.004936849],"category_scores_gemma":[0.015052283,0.00045274897,0.00070754794,0.0007162839,0.0006139665,0.0031052849,0.0019266679,0.0026549404,0.0021663804],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083124713,0.000526133,0.004244968,0.00060940423,0.00008783413,0.00028901917,0.0007076255,0.0871178,0.035632037,0.010770745,0.016885227,0.842298],"study_design_scores_gemma":[0.00021480258,0.00044459853,0.0014186973,0.000046547266,0.000060619408,0.0001887872,0.00014575814,0.93121415,0.019580852,0.03757767,0.009060616,0.000046819958],"about_ca_topic_score_codex":0.00072432373,"about_ca_topic_score_gemma":0.0013506841,"teacher_disagreement_score":0.004936849,"about_ca_system_score_codex":0.00063257327,"about_ca_system_score_gemma":0.0010291649,"threshold_uncertainty_score":0.016515374},"labels":[],"label_agreement":null},{"id":"W6958058837","doi":"10.60692/mfwkg-4eb72","title":"QRelScore: Better Evaluating Generated Questions with Deeper Understanding of Context-aware Relevance","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Canadian Institute for Advanced Research","funders":"","keywords":"Relevance (law); Adversarial system; Matching (statistics); Context (archaeology); Metric (unit); Rendering (computer graphics)","score_opus":0.0923810007886917,"score_gpt":0.2493279509683739,"score_spread":0.1569469501796822,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958058837","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2880936,0.0068323235,0.6541147,0.00083590474,0.0006039947,0.0017380695,0.005483459,0.025998915,0.016299004],"genre_scores_gemma":[0.7266391,0.00045560216,0.25982156,0.00039748638,0.00015355054,0.0005592172,0.007556923,0.00079986453,0.003616724],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9920644,0.003794085,0.0005638251,0.0014416118,0.001900898,0.00023517938],"domain_scores_gemma":[0.9729672,0.020248624,0.001516998,0.002320073,0.002304799,0.0006423013],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007915342,0.0018750947,0.0010500501,0.0036919268,0.00049895706,0.0022497494,0.001237513,0.002021905,0.0047504315],"category_scores_gemma":[0.053257767,0.0002376516,0.0007196791,0.0011866089,0.0006821225,0.003306251,0.0022975435,0.0013971746,0.0015212442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021728494,0.00080960407,0.0286116,0.002926109,0.0006054227,0.0003839027,0.0015523841,0.083132915,0.04982031,0.009973132,0.031370748,0.7886411],"study_design_scores_gemma":[0.00026092955,0.0016982985,0.02630575,0.00022232215,0.00026373527,0.0007093057,0.0005152348,0.8778665,0.050414883,0.018730145,0.022753114,0.00025981042],"about_ca_topic_score_codex":0.002213994,"about_ca_topic_score_gemma":0.003573897,"teacher_disagreement_score":0.007915342,"about_ca_system_score_codex":0.00093896996,"about_ca_system_score_gemma":0.0011694696,"threshold_uncertainty_score":0.04186088},"labels":[],"label_agreement":null},{"id":"W6958059733","doi":"10.60692/9td5k-b0448","title":"Integrating Semantics and Neighborhood Information with Graph-Driven Generative Models for Document Retrieval","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Semantics (computer science); Generative grammar; Computational linguistics; Joint (building); Natural language; Natural (archaeology); Computational semantics; Natural language generation","score_opus":0.026819208894033702,"score_gpt":0.21204143457446367,"score_spread":0.18522222568042998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958059733","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030943898,0.0027331982,0.9616676,0.00064017286,0.00010110727,0.000101250844,0.0004953208,0.0016615322,0.0016558392],"genre_scores_gemma":[0.75909466,0.0020869158,0.22902003,0.0004248831,0.00037823597,0.0003763448,0.0021665904,0.00078104966,0.0056711864],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921596,0.00037299973,0.000047985985,0.00018012266,0.000114684495,0.00006820301],"domain_scores_gemma":[0.9972154,0.0021061737,0.00014333546,0.00023146071,0.00020602312,0.00009759837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018062047,0.00087264227,0.0019321295,0.0027795646,0.00076894666,0.0015865074,0.00232926,0.0016349279,0.0022487605],"category_scores_gemma":[0.005976652,0.0009700301,0.0020531095,0.0025728755,0.00089725776,0.0039107515,0.0017166603,0.0016422895,0.0012173877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048461807,0.00026204638,0.0038292452,0.00025725664,0.00035942422,0.00018869994,0.00038147075,0.72729236,0.0025119283,0.06984085,0.009568892,0.1850232],"study_design_scores_gemma":[0.000017039383,0.000017099108,0.00013688572,0.000005490362,0.000028027442,0.00002138043,0.000012480958,0.9687205,0.0001534589,0.030386293,0.0004926002,0.000008883718],"about_ca_topic_score_codex":0.016957354,"about_ca_topic_score_gemma":0.029795527,"teacher_disagreement_score":0.016957354,"about_ca_system_score_codex":0.0014250741,"about_ca_system_score_gemma":0.001130361,"threshold_uncertainty_score":0.033717275},"labels":[],"label_agreement":null},{"id":"W6958088051","doi":"10.60692/2dzhz-dpc57","title":"On-the-Fly Attention Modulation for Neural Generation","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Modulation (music); Artificial neural network; Neural activity; Prior probability; Computation; Deep neural networks; Simple (philosophy)","score_opus":0.07389547561211676,"score_gpt":0.2246698726626971,"score_spread":0.15077439705058035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958088051","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13565348,0.00137626,0.8435111,0.00060142466,0.00020301345,0.00014432485,0.00030580373,0.010603523,0.0076011214],"genre_scores_gemma":[0.8767486,0.00024410513,0.11946398,0.00021237224,0.00006821641,0.00012555439,0.00024672944,0.00029389653,0.00259649],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999762,0.000082226245,0.000014879262,0.00006685494,0.000046680198,0.000027276294],"domain_scores_gemma":[0.999303,0.00042857527,0.00005254633,0.00011962153,0.000064256776,0.00003204416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00046341636,0.0005631474,0.00031649633,0.00029104634,0.00019741303,0.00047146843,0.0010310615,0.00059145154,0.0034576517],"category_scores_gemma":[0.0027442428,0.00020160127,0.0002680534,0.00026616023,0.0002987478,0.0010422688,0.0008277707,0.00097604486,0.00065711985],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003131087,0.0002213983,0.0013631383,0.00020129533,0.000067003966,0.0001637564,0.00021832116,0.10136976,0.08451644,0.008542587,0.0048594647,0.79816353],"study_design_scores_gemma":[0.00003705117,0.00008516411,0.00069081323,0.000012402897,0.000022481972,0.00004641287,0.000023781186,0.95891595,0.028739328,0.009107614,0.002308728,0.000010280303],"about_ca_topic_score_codex":0.0018407282,"about_ca_topic_score_gemma":0.0034657957,"teacher_disagreement_score":0.0034576517,"about_ca_system_score_codex":0.0005566823,"about_ca_system_score_gemma":0.00047147044,"threshold_uncertainty_score":0.011567056},"labels":[],"label_agreement":null},{"id":"W6958176606","doi":"10.6084/m9.figshare.18666709.v1","title":"Additional file 2 of Adapting systematic scoping study methods to identify cancer-specific physical activity opportunities in Ontario, Canada","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Queen's University; University of Toronto","funders":"","keywords":"Physical activity; Key (lock); Data collection; Identification (biology); Activity recognition","score_opus":0.24519623718171615,"score_gpt":0.36165878692643905,"score_spread":0.11646254974472289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958176606","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00013212868,0.000067591325,0.00014066622,0.00009463981,0.000015336063,0.0006260743,0.9974189,0.00010779968,0.001396832],"genre_scores_gemma":[0.011737282,0.0011818893,0.012010389,0.00066731893,0.000081307095,0.037038483,0.9122766,0.0007327799,0.02427396],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9975738,0.00039568768,0.00088790915,0.00033020406,0.0005135076,0.00029900292],"domain_scores_gemma":[0.93731034,0.04331961,0.0035765425,0.0018210001,0.0130578065,0.0009147011],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.005119256,0.0013123389,0.0022860668,0.010994059,0.0020406526,0.0026413752,0.0021447907,0.0013072671,0.77314717],"category_scores_gemma":[0.07015838,0.0012057091,0.0028815262,0.020085504,0.00057358027,0.0021361317,0.0019510104,0.0007898882,0.03801677],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005298816,0.000052470292,0.0031727822,0.06511729,0.00021866521,0.00008221588,0.0005624789,0.00058225065,0.00013748376,0.0014519677,0.9111387,0.016953787],"study_design_scores_gemma":[0.008550278,0.0001810887,0.048753615,0.04909089,0.0014521823,0.00019786032,0.0013126276,0.0012445113,0.00045261488,0.004508256,0.8840403,0.00021583177],"about_ca_topic_score_codex":0.44501457,"about_ca_topic_score_gemma":0.70041674,"teacher_disagreement_score":0.77314717,"about_ca_system_score_codex":0.010783036,"about_ca_system_score_gemma":0.03169283,"threshold_uncertainty_score":0.8848486},"labels":[],"label_agreement":null},{"id":"W6958286424","doi":"10.60692/vq1jq-mj682","title":"Bridging the Gap between Language Models and Cross-Lingual Sequence Labeling","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Bridging (networking); Language model; Sequence labeling; Margin (machine learning); Regularization (linguistics); Structured prediction; Task (project management); Consistency (knowledge bases); Sequence (biology); Security token","score_opus":0.08578039782984098,"score_gpt":0.2706474233342497,"score_spread":0.1848670255044087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958286424","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06938198,0.0022273352,0.9063878,0.0007892892,0.00026083842,0.00020170453,0.000923997,0.015558154,0.0042689247],"genre_scores_gemma":[0.52794147,0.00076519774,0.4499929,0.0015769161,0.00019269697,0.0006186778,0.008218139,0.0021870574,0.008506957],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99627835,0.0014006875,0.00017524256,0.0016357997,0.0003227258,0.00018719699],"domain_scores_gemma":[0.9927516,0.004076529,0.000324751,0.00183409,0.00079037633,0.00022257505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035659687,0.0023920587,0.0015503583,0.0011940054,0.0009780863,0.0018637219,0.003580983,0.0022522255,0.0034600906],"category_scores_gemma":[0.01163255,0.0011416228,0.0012584248,0.0013504925,0.0014222026,0.003987081,0.0037612338,0.005003884,0.0033706115],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060256454,0.0004189823,0.0041126283,0.0005856162,0.00031474658,0.00037562614,0.0010085126,0.15635978,0.02556601,0.007434434,0.0156056015,0.78761554],"study_design_scores_gemma":[0.000038936698,0.00017600319,0.0014038148,0.000075529075,0.00006939251,0.0001855675,0.0002972889,0.9609789,0.015673393,0.013711118,0.0073373355,0.000052674364],"about_ca_topic_score_codex":0.007320263,"about_ca_topic_score_gemma":0.013128566,"teacher_disagreement_score":0.007320263,"about_ca_system_score_codex":0.0013539484,"about_ca_system_score_gemma":0.0019080671,"threshold_uncertainty_score":0.01885885},"labels":[],"label_agreement":null},{"id":"W6958462973","doi":"10.60692/bzw3c-ctp50","title":"HotpotQA: A Dataset for Diverse, Explainable Multi-hop Question Answering","year":2018,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Question answering; Natural language; Natural (archaeology); Empirical research","score_opus":0.06763872202132244,"score_gpt":0.2565929837434248,"score_spread":0.1889542617221024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958462973","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02726297,0.0036879466,0.009488702,0.001967398,0.000407583,0.00075037545,0.93725955,0.01299933,0.0061762217],"genre_scores_gemma":[0.020112354,0.00025531423,0.011510955,0.00035073952,0.000062635925,0.00046639823,0.9647473,0.00020231873,0.0022920624],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977349,0.0006412258,0.0002909618,0.0006221236,0.00050304574,0.00020783514],"domain_scores_gemma":[0.9955974,0.0019265418,0.00024192125,0.0009414904,0.0008599926,0.00043268222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019548899,0.0029111372,0.0017574506,0.0043620984,0.0018660129,0.0018580775,0.00438606,0.004959098,0.014736331],"category_scores_gemma":[0.0106695965,0.0006748308,0.0019975158,0.0033420348,0.00071250746,0.003678318,0.0039921016,0.002338303,0.012106093],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062670134,0.00047080347,0.005840034,0.0022055407,0.0002624864,0.00049896346,0.0004777852,0.0027635372,0.0033908184,0.002308218,0.9414163,0.039738838],"study_design_scores_gemma":[0.0016805865,0.0006213274,0.03694228,0.00078256556,0.0004335794,0.0013596402,0.002565175,0.07901051,0.008418934,0.018399037,0.8494439,0.00034246396],"about_ca_topic_score_codex":0.03201551,"about_ca_topic_score_gemma":0.06417438,"teacher_disagreement_score":0.03201551,"about_ca_system_score_codex":0.0018490859,"about_ca_system_score_gemma":0.002356967,"threshold_uncertainty_score":0.0636583},"labels":[],"label_agreement":null},{"id":"W6958788410","doi":"10.60841/000000005","title":"Data from: Unveiling What Makes Saturn Ring: Quantifying the Amplitudes of Saturn's Planetary Normal-Mode Oscillations and Trends in C-Ring Properties Using Kronoseismology (VII)","year":2024,"lang":"en","type":"dataset","venue":"University of Idaho | Library","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Theoretical Astrophysics","funders":"National Aeronautics and Space Administration","keywords":"Amplitude; Saturn; Occultation; Satellite; Planet; Gravitational wave; Rotation (mathematics); Atmospheric wave","score_opus":0.10088013322198112,"score_gpt":0.26492482395143796,"score_spread":0.16404469072945682,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958788410","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002210834,0.0003067653,0.00007991423,0.000162483,0.000032697702,0.000013718927,0.9961331,0.00038769815,0.0006727678],"genre_scores_gemma":[0.0023520298,0.00009310543,0.00028812053,0.000033835942,0.0000106109765,0.00004307523,0.996765,0.00002375802,0.00039035155],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99935883,0.00010452315,0.0000735328,0.00018821991,0.00015568527,0.00011923947],"domain_scores_gemma":[0.9988796,0.00031260637,0.00013637837,0.00024128742,0.00028307832,0.00014705474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008584269,0.0016784758,0.0009703671,0.0030127156,0.00084483396,0.0015933827,0.001755891,0.0018076736,0.011019997],"category_scores_gemma":[0.0035256275,0.0003460472,0.0011882783,0.004394176,0.00035364233,0.00091089105,0.0014917176,0.001090685,0.018270599],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019460854,0.00008104016,0.013673019,0.0017540068,0.00015490486,0.000095088435,0.000116092306,0.0012636317,0.0005835526,0.00051326264,0.9731557,0.008415006],"study_design_scores_gemma":[0.00055743195,0.00006889023,0.071121514,0.00046603061,0.00014052789,0.00021095373,0.00044178008,0.0037286254,0.0014453377,0.0011488716,0.9205918,0.00007817058],"about_ca_topic_score_codex":0.042330325,"about_ca_topic_score_gemma":0.069785126,"teacher_disagreement_score":0.042330325,"about_ca_system_score_codex":0.0013069168,"about_ca_system_score_gemma":0.0016528197,"threshold_uncertainty_score":0.0841679},"labels":[],"label_agreement":null},{"id":"W6958885860","doi":"10.6084/m9.figshare.c.7050923","title":"Canadians’ knowledge of cancer risk factors and belief in cancer myths","year":2024,"lang":"en","type":"other","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Centre for Disease Control; Canadian Cancer Society; University of Calgary; University of British Columbia","funders":"","keywords":"Cancer; Feeling; Cognition; Preference; Population; Causality (physics); Risk perception; Risk factor","score_opus":0.038872424284382597,"score_gpt":0.2953458920857877,"score_spread":0.2564734678014051,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958885860","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98751384,0.0011204518,0.00006850269,0.0014806809,0.000020918116,0.000024728675,0.0008853464,0.000005394735,0.008880036],"genre_scores_gemma":[0.99863476,0.00059408403,0.000065924876,0.000100735415,0.0000067021015,0.0000047475146,0.00019740613,0.0000010362065,0.00039462745],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99895763,0.00011486751,0.00005687026,0.000081078535,0.00054213253,0.00024746655],"domain_scores_gemma":[0.9950571,0.0009223891,0.0013164824,0.00010327406,0.0017286435,0.0008721427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014223299,0.00016212834,0.00023916754,0.0017744781,0.0023758893,0.001334931,0.00058419845,0.00044493563,0.005460274],"category_scores_gemma":[0.008180966,0.00016277691,0.00041135153,0.0022070506,0.001212816,0.00046521792,0.0007455706,0.00072461093,0.00016996231],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000082687795,0.00005952082,0.95802796,0.00019762434,0.000052945274,0.00017309435,0.017160563,0.00009821751,0.00017647768,0.0004029466,0.0023196016,0.021248229],"study_design_scores_gemma":[0.000008855464,0.00005451015,0.9755558,0.00021032989,0.00003077563,0.0001419199,0.019534674,0.00018602822,0.000059746766,0.00015115293,0.0040382803,0.000027895012],"about_ca_topic_score_codex":0.9563534,"about_ca_topic_score_gemma":0.95081145,"teacher_disagreement_score":0.043646574,"about_ca_system_score_codex":0.007971891,"about_ca_system_score_gemma":0.013526193,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W6959768449","doi":"10.1051/0004-6361/202450701/pdf","title":"Validation of the RR Lyrae period determination in the Pan-STARRS PS1 3","year":2024,"lang":"en","type":"article","venue":"Springer Link (Chiba Institute of Technology)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Los Alamos National Laboratory; Planetary Science Division; Science Mission Directorate; Smithsonian Astrophysical Observatory; Max-Planck-Institut für Astronomie; Eötvös Loránd Tudományegyetem; Magyar Tudományos Akadémia; Space Telescope Science Institute; National Central University; Gordon and Betty Moore Foundation; Queen's University; Nemzeti Kutatási Fejlesztési és Innovációs Hivatal; Johns Hopkins University; Queen's University Belfast; National Science Foundation; European Space Agency; National Aeronautics and Space Administration; Durham University; Smithsonian Institution","keywords":"RR Lyrae variable; Light curve; Stars; Variable star; Photometry (optics); Ecliptic","score_opus":0.01732327968529611,"score_gpt":0.24973936295080107,"score_spread":0.23241608326550495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6959768449","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97172666,0.0006391536,0.013756689,0.000051633357,0.000068892194,0.00010257146,0.006898979,0.00059670454,0.0061586355],"genre_scores_gemma":[0.97099394,0.00014312784,0.010638003,0.00006285981,0.000048496353,0.000070699054,0.017032366,0.00020684546,0.00080361107],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975647,0.00044096116,0.00025171306,0.00093476556,0.00060584885,0.00020202155],"domain_scores_gemma":[0.99476045,0.0006694823,0.0012182504,0.0014982284,0.0014768712,0.0003767625],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00352612,0.00041986338,0.000314945,0.0020612197,0.0004038165,0.0009109023,0.00052942015,0.00031617607,0.0015128562],"category_scores_gemma":[0.005151406,0.00019673564,0.00045985178,0.001448763,0.00026674833,0.0006124122,0.0010409222,0.00023222472,0.0016534375],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048900157,0.00008082762,0.899425,0.0002892159,0.00023932749,0.00018564072,0.00053905154,0.0034130374,0.033934295,0.00048523338,0.0025773405,0.058341943],"study_design_scores_gemma":[0.000020576495,0.00012918023,0.9806805,0.000036856196,0.0000647258,0.00022181518,0.0001746099,0.0038532265,0.0056108786,0.00009865779,0.009100192,0.000008695305],"about_ca_topic_score_codex":0.0040700138,"about_ca_topic_score_gemma":0.004874261,"teacher_disagreement_score":0.0040700138,"about_ca_system_score_codex":0.00039677846,"about_ca_system_score_gemma":0.0004891873,"threshold_uncertainty_score":0.018648148},"labels":[],"label_agreement":null},{"id":"W6966644381","doi":"10.3897/zookeys.75.767.figures14-19","title":"Figures 14-19 from: Brunke A, Marshall S (2011) Contributions to the faunistics and bionomics of Staphylinidae (Coleoptera) in northeastern North America: discoveries made through study of the University of Guelph Insect Collection, Ontario, Canada. ZooKeys 75: 29-68. https://doi.org/10.3897/zookeys.75.767","year":2011,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Aedeagus; Bionomics; Insect; Larva","score_opus":0.028661783170707132,"score_gpt":0.209284761841337,"score_spread":0.18062297867062987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6966644381","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006208205,0.0009668887,0.0017592305,0.00070557644,0.001091357,0.00018905476,0.19635215,0.0019822984,0.79633266],"genre_scores_gemma":[0.007063719,0.0023287486,0.0053382074,0.00024825372,0.00028157295,0.00025694983,0.15457182,0.0022163393,0.8276945],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996207,0.000029244939,0.000022762213,0.00005594351,0.00020861397,0.00006269524],"domain_scores_gemma":[0.9986411,0.0002492538,0.000092712966,0.00013797774,0.0007085136,0.0001704412],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00032992996,0.0010597238,0.00057320483,0.0027786277,0.0011070268,0.0023248505,0.0010241639,0.0006016937,0.80843747],"category_scores_gemma":[0.003922149,0.00036304403,0.00065317564,0.0072079515,0.0005340936,0.0024669268,0.0010764105,0.0009564962,0.6055675],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008194424,0.0000033950416,0.00012442606,0.00007377395,0.0000015139768,0.000014731997,0.000054855416,0.00006387789,0.000031209383,0.00083643966,0.9896763,0.009111178],"study_design_scores_gemma":[0.0000033939791,0.0000017251137,0.00067571533,0.000063156134,0.0000016902053,0.000030540337,0.000057781206,0.000042517233,0.00003251672,0.0003787141,0.9987099,0.0000022981462],"about_ca_topic_score_codex":0.0754578,"about_ca_topic_score_gemma":0.15157564,"teacher_disagreement_score":0.9245422,"about_ca_system_score_codex":0.0023498177,"about_ca_system_score_gemma":0.0025235564,"threshold_uncertainty_score":0.27324063},"labels":[],"label_agreement":null},{"id":"W6966665280","doi":"10.48448/egen-6z71","title":"Responsible AI Considerations in Text Summarization Research: A Review of Current Practices","year":2023,"lang":"en","type":"other","venue":"Open MIND","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Current (fluid); Term (time); Feature (linguistics); Key (lock)","score_opus":0.48856608543874946,"score_gpt":0.515350177369354,"score_spread":0.026784091930604492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6966665280","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012346129,0.8721527,0.07553731,0.025269715,0.0013700898,0.00016345375,0.00019746808,0.00051576307,0.023558868],"genre_scores_gemma":[0.032222837,0.83008766,0.12043367,0.0066975662,0.0047245007,0.00034568843,0.0005752242,0.0003503933,0.0045625283],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99041736,0.0039590155,0.0012243715,0.0012843001,0.0028931468,0.00022184513],"domain_scores_gemma":[0.8399387,0.14263126,0.0030957381,0.0038915742,0.009240528,0.0012021362],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.025735693,0.0009899428,0.0017752665,0.010624666,0.001844747,0.00963195,0.0038026874,0.003727022,0.010363176],"category_scores_gemma":[0.058464408,0.0008116424,0.0011186209,0.010972563,0.0049951435,0.015603923,0.0036870595,0.004122316,0.0040443474],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007410139,0.00011002509,0.0006159215,0.011776011,0.00009815622,0.000069750495,0.0012796748,0.00062096753,0.0006930561,0.05003083,0.018513365,0.91611814],"study_design_scores_gemma":[0.00006281712,0.00011571468,0.0020768445,0.02320245,0.00036411086,0.0005427906,0.0035862506,0.006856987,0.0024614278,0.28810617,0.67248493,0.00013950578],"about_ca_topic_score_codex":0.0036048135,"about_ca_topic_score_gemma":0.0060520982,"teacher_disagreement_score":0.9742643,"about_ca_system_score_codex":0.0026753792,"about_ca_system_score_gemma":0.004972613,"threshold_uncertainty_score":0.13610494},"labels":[],"label_agreement":null},{"id":"W6966841497","doi":"10.48448/2s2c-3s80","title":"Context-aware Adversarial Training for Name Regularity Bias in Named Entity Recognition","year":2021,"lang":"en","type":"other","venue":"Underline Science Inc.","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Université de Montréal","funders":"","keywords":"Adversarial system; Training (meteorology); Named-entity recognition; Training set; Pattern recognition (psychology); Key (lock)","score_opus":0.15597899223028241,"score_gpt":0.3242821735167069,"score_spread":0.1683031812864245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6966841497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043172073,0.0016336696,0.9459804,0.0008248799,0.00029326725,0.000046535497,0.00056307664,0.0029115535,0.0045745643],"genre_scores_gemma":[0.8429946,0.0008952439,0.13898459,0.0006784155,0.0004734991,0.00009815559,0.0028746205,0.0004861986,0.012514661],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992107,0.00033960605,0.000033503213,0.00023956924,0.00009850325,0.00007813419],"domain_scores_gemma":[0.9975677,0.0015598064,0.00012902098,0.00048695583,0.00018817287,0.00006834038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015024801,0.00076869794,0.00076380384,0.0004783467,0.0004475208,0.00068647915,0.0013469594,0.00092185294,0.0036739158],"category_scores_gemma":[0.005240375,0.0003279743,0.00045622262,0.00079718215,0.00056952395,0.0015558292,0.00177192,0.0022570514,0.0020930285],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000536289,0.00020340196,0.0027866336,0.00015794198,0.0001397468,0.00016917898,0.00015998154,0.41017577,0.010112736,0.020262372,0.026069727,0.5292262],"study_design_scores_gemma":[0.000006616984,0.000022923554,0.000283702,0.0000100842035,0.000011890087,0.000031411586,0.000011082442,0.98883444,0.0021490226,0.0074238963,0.0012081896,0.0000068258973],"about_ca_topic_score_codex":0.002966564,"about_ca_topic_score_gemma":0.004680465,"teacher_disagreement_score":0.0036739158,"about_ca_system_score_codex":0.00040906604,"about_ca_system_score_gemma":0.000651742,"threshold_uncertainty_score":0.012290478},"labels":[],"label_agreement":null},{"id":"W6968829696","doi":"10.5281/zenodo.4265632","title":"MeDAL: Medical Abbreviation Disambiguation Dataset for Natural Language Understanding Pretraining","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Download; Upload; Natural language; Text messaging; Natural (archaeology)","score_opus":0.11244202040359885,"score_gpt":0.2983701788779148,"score_spread":0.18592815847431593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6968829696","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009205873,0.0013088614,0.0062352396,0.0015052364,0.00046674497,0.0007342814,0.96557873,0.008515803,0.0064491974],"genre_scores_gemma":[0.006547213,0.00016659073,0.010000189,0.0005370414,0.000050767685,0.0006177048,0.9800149,0.00025699622,0.0018085919],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975963,0.0005512353,0.0004334199,0.0007936361,0.00044935665,0.00017586944],"domain_scores_gemma":[0.99693346,0.0010795224,0.00022901174,0.0006465875,0.0007285353,0.00038294398],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002098877,0.0023568526,0.0012190691,0.0040993392,0.0014903743,0.0013438668,0.0030184903,0.0029825794,0.02802622],"category_scores_gemma":[0.009784666,0.00052431016,0.0016058291,0.0023972017,0.00080977904,0.0017716492,0.0026457945,0.0026198905,0.038036942],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041054623,0.00014790703,0.0022641586,0.0014288996,0.00008148145,0.00034416615,0.00017172692,0.0009887217,0.0023761739,0.000997667,0.9595207,0.031267934],"study_design_scores_gemma":[0.0010723949,0.00036019148,0.018380538,0.00059435196,0.00017010153,0.002118808,0.000631833,0.012619841,0.009988935,0.0046029515,0.9492639,0.00019624026],"about_ca_topic_score_codex":0.013416283,"about_ca_topic_score_gemma":0.029274162,"teacher_disagreement_score":0.02802622,"about_ca_system_score_codex":0.0019459457,"about_ca_system_score_gemma":0.003625447,"threshold_uncertainty_score":0.09375703},"labels":[],"label_agreement":null},{"id":"W6969110423","doi":"10.5281/zenodo.16875905","title":"Agentic and Non-Agentic Multi-Hop Systems for Medical Question Answering","year":2025,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Task (project management); Pipeline (software); Question answering; Questions and answers; Joint (building); Semantics (computer science); Interrogative word","score_opus":0.030864254415028207,"score_gpt":0.2675632992729495,"score_spread":0.23669904485792126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6969110423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026593411,0.0007563605,0.9168474,0.0016738338,0.00019582102,0.0006886756,0.0012517159,0.04601718,0.005975682],"genre_scores_gemma":[0.26878646,0.00025249124,0.71758336,0.0007509224,0.00011941981,0.0005756232,0.003687493,0.0012758733,0.0069684135],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99682057,0.0014088497,0.00025652093,0.0007390459,0.0006047189,0.00017024101],"domain_scores_gemma":[0.9935854,0.0035225253,0.00026538293,0.001502982,0.00074434094,0.00037932213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005027283,0.0008877372,0.0008799285,0.0014364271,0.0009886628,0.0029284733,0.0032548,0.002258306,0.008285554],"category_scores_gemma":[0.011575791,0.00063811924,0.0011400572,0.00092162547,0.0009904999,0.0039288313,0.0051847007,0.0019427635,0.003496507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028732365,0.0014310473,0.0066038948,0.001800051,0.00076010113,0.0010860257,0.0026180283,0.121794775,0.050662905,0.066248804,0.07640118,0.6677199],"study_design_scores_gemma":[0.00019004644,0.00015047808,0.00062564696,0.000033273664,0.000064189844,0.00011657738,0.0002376036,0.92544526,0.017025072,0.027325802,0.028730832,0.000055168275],"about_ca_topic_score_codex":0.006256642,"about_ca_topic_score_gemma":0.008570307,"teacher_disagreement_score":0.008285554,"about_ca_system_score_codex":0.0012114714,"about_ca_system_score_gemma":0.002017825,"threshold_uncertainty_score":0.027717948},"labels":[],"label_agreement":null},{"id":"W6976533378","doi":"10.60692/ecg0e-c5z35","title":"Balaur: Language Model Pretraining with Lexical Semantic Relations","year":2023,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Generalization; Inference; Set (abstract data type); Meaning (existential); Semantics (computer science); Lexical semantics; Interface (matter); Language model","score_opus":0.041358990059133474,"score_gpt":0.22357447574456035,"score_spread":0.1822154856854269,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976533378","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03291513,0.00071476336,0.9203281,0.00071395264,0.0003291596,0.00018126886,0.0012164867,0.038949866,0.00465132],"genre_scores_gemma":[0.45757905,0.0003802234,0.5176889,0.0012647255,0.000197857,0.0006714777,0.006557065,0.0033269464,0.012333794],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99947745,0.00016785506,0.000028150607,0.0002112635,0.00006345181,0.00005179768],"domain_scores_gemma":[0.9982577,0.0011886106,0.000052608986,0.00026156803,0.00016760577,0.00007194006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013195217,0.0014358672,0.00086520775,0.0006893537,0.00048771207,0.0011831158,0.0023466684,0.0013554881,0.011134171],"category_scores_gemma":[0.005147214,0.0007857011,0.0012405271,0.00059983763,0.0005843591,0.0029145053,0.0021160357,0.00454892,0.0048312233],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000667738,0.0003803659,0.0026468642,0.00044960377,0.0003650934,0.00024546677,0.0005510563,0.18963858,0.021250391,0.01215148,0.04243851,0.72921485],"study_design_scores_gemma":[0.000048640417,0.00006198304,0.00036398336,0.00002608721,0.00003565855,0.000050116512,0.000070746355,0.97876567,0.005468912,0.011092511,0.003992921,0.000022806526],"about_ca_topic_score_codex":0.006038872,"about_ca_topic_score_gemma":0.015268861,"teacher_disagreement_score":0.011134171,"about_ca_system_score_codex":0.00073993637,"about_ca_system_score_gemma":0.0011490849,"threshold_uncertainty_score":0.03724754},"labels":[],"label_agreement":null},{"id":"W6976597069","doi":"10.60692/ck6mt-jht14","title":"f-Divergence Minimization for Sequence-Level Knowledge Distillation","year":2023,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Advanced Research; McGill University; University of Alberta","funders":"","keywords":"Distillation; Divergence (linguistics); Minification; Process (computing); Decomposition; Natural language","score_opus":0.15298446752216682,"score_gpt":0.27007000950888016,"score_spread":0.11708554198671334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976597069","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010000018,0.00043967617,0.98728997,0.00038414574,0.00003643658,0.000046487643,0.00014395619,0.00048360412,0.0011756782],"genre_scores_gemma":[0.4101383,0.00076160335,0.578282,0.000734774,0.00020580279,0.00036708036,0.0017492345,0.00048197398,0.0072791623],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982168,0.00061425805,0.00012513828,0.00043608833,0.0004452262,0.00016238543],"domain_scores_gemma":[0.9962081,0.0025810492,0.00018361885,0.00046442053,0.00040790645,0.00015488482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038276985,0.0012199809,0.0017623319,0.0013689448,0.0007930717,0.0015761964,0.0024215023,0.002311929,0.0029101148],"category_scores_gemma":[0.010994599,0.00043793436,0.0010525109,0.0017324035,0.0018586586,0.003895056,0.0026222228,0.0035540303,0.0011036793],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001721361,0.00020439603,0.0011430101,0.00026194958,0.00008755472,0.00011584501,0.00017817438,0.64963967,0.003047674,0.0684181,0.008015264,0.26871625],"study_design_scores_gemma":[0.000010767847,0.000033303313,0.0001148876,0.000012373049,0.0000061894134,0.000030501087,0.000016421887,0.96221524,0.0010645515,0.03520898,0.0012748756,0.0000118338385],"about_ca_topic_score_codex":0.004251463,"about_ca_topic_score_gemma":0.0051350803,"teacher_disagreement_score":0.004251463,"about_ca_system_score_codex":0.0021375765,"about_ca_system_score_gemma":0.0025518092,"threshold_uncertainty_score":0.020243108},"labels":[],"label_agreement":null},{"id":"W6976617253","doi":"10.60692/e8xhv-9km80","title":"HotpotQA: A Dataset for Diverse, Explainable Multi-hop Question Answering","year":2018,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Question answering; Natural language; Natural (archaeology); Empirical research","score_opus":0.06763872202132244,"score_gpt":0.2565929837434248,"score_spread":0.1889542617221024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976617253","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02726297,0.0036879466,0.009488702,0.001967398,0.000407583,0.00075037545,0.93725955,0.01299933,0.0061762217],"genre_scores_gemma":[0.020112354,0.00025531423,0.011510955,0.00035073952,0.000062635925,0.00046639823,0.9647473,0.00020231873,0.0022920624],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977349,0.0006412258,0.0002909618,0.0006221236,0.00050304574,0.00020783514],"domain_scores_gemma":[0.9955974,0.0019265418,0.00024192125,0.0009414904,0.0008599926,0.00043268222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019548899,0.0029111372,0.0017574506,0.0043620984,0.0018660129,0.0018580775,0.00438606,0.004959098,0.014736331],"category_scores_gemma":[0.0106695965,0.0006748308,0.0019975158,0.0033420348,0.00071250746,0.003678318,0.0039921016,0.002338303,0.012106093],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062670134,0.00047080347,0.005840034,0.0022055407,0.0002624864,0.00049896346,0.0004777852,0.0027635372,0.0033908184,0.002308218,0.9414163,0.039738838],"study_design_scores_gemma":[0.0016805865,0.0006213274,0.03694228,0.00078256556,0.0004335794,0.0013596402,0.002565175,0.07901051,0.008418934,0.018399037,0.8494439,0.00034246396],"about_ca_topic_score_codex":0.03201551,"about_ca_topic_score_gemma":0.06417438,"teacher_disagreement_score":0.03201551,"about_ca_system_score_codex":0.0018490859,"about_ca_system_score_gemma":0.002356967,"threshold_uncertainty_score":0.0636583},"labels":[],"label_agreement":null},{"id":"W6976775207","doi":"10.60692/rwdwm-ste41","title":"MasakhaNER 2.0: Africa-centric Transfer Learning for Named Entity Recognition","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Mila - Quebec Artificial Intelligence Institute","funders":"","keywords":"Named-entity recognition; Task (project management); Transfer of learning; Benchmark (surveying); Transfer (computing); Cover (algebra); Training set; Sequence labeling","score_opus":0.05388125223713194,"score_gpt":0.20273617661652574,"score_spread":0.1488549243793938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976775207","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09970919,0.003899231,0.6484873,0.0016199832,0.0009802385,0.00141338,0.028839024,0.20121934,0.013832286],"genre_scores_gemma":[0.31667835,0.0017846858,0.5399263,0.0010286596,0.00026127545,0.0022058138,0.11057944,0.0073250867,0.020210452],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99852437,0.0006142534,0.00008545979,0.00043388145,0.0001903313,0.0001517549],"domain_scores_gemma":[0.9985403,0.00058000494,0.0000836354,0.0004828894,0.00022830754,0.00008481348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004589529,0.0023051768,0.0011371559,0.0024380556,0.0012849284,0.0017824731,0.002897191,0.0016178765,0.008176409],"category_scores_gemma":[0.0064952564,0.0007390931,0.0015033847,0.0015570193,0.0004869811,0.0041645234,0.004606307,0.002595765,0.006973213],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015873738,0.0010674186,0.010889545,0.0012971611,0.0011359248,0.0008845027,0.0012809231,0.057644628,0.02155802,0.010807059,0.1933499,0.6984976],"study_design_scores_gemma":[0.00037114407,0.00068335095,0.008859975,0.0002448982,0.00031633905,0.0012791678,0.00083854486,0.7551333,0.051530354,0.024807395,0.1557137,0.00022181771],"about_ca_topic_score_codex":0.0060476055,"about_ca_topic_score_gemma":0.011798946,"teacher_disagreement_score":0.008176409,"about_ca_system_score_codex":0.00086040696,"about_ca_system_score_gemma":0.0017190123,"threshold_uncertainty_score":0.02735281},"labels":[],"label_agreement":null},{"id":"W6976831743","doi":"10.60692/2m81x-g0y08","title":"Feeding What You Need by Understanding What You Learned","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Pipeline (software); Interpretation (philosophy); Comprehension; Training set; Reading (process); Curriculum; Deep learning","score_opus":0.08032085247982688,"score_gpt":0.22488749781583378,"score_spread":0.14456664533600688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976831743","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23133262,0.0011440978,0.72229993,0.005410756,0.00031911177,0.00060220214,0.004202355,0.01239589,0.022292992],"genre_scores_gemma":[0.5394021,0.0011280014,0.43843746,0.00087740034,0.00012177066,0.0006084441,0.007920262,0.00070227374,0.010802225],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990031,0.00037138988,0.000056486366,0.00036083127,0.00014982959,0.00005843822],"domain_scores_gemma":[0.99446243,0.0037743873,0.0002574386,0.000761646,0.00054599636,0.00019797451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015534571,0.001514397,0.00081097323,0.001229219,0.0005218633,0.0024311198,0.0013875914,0.0018384807,0.013536249],"category_scores_gemma":[0.018382104,0.00059620856,0.001135781,0.0010862863,0.0004768337,0.0068519562,0.002497848,0.002620915,0.005926637],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004591481,0.0010010756,0.024561157,0.0014882841,0.00017702261,0.00043166676,0.0038107017,0.03701868,0.030689504,0.014761017,0.032484826,0.85311687],"study_design_scores_gemma":[0.00011616268,0.00085884763,0.016874518,0.00058307115,0.00026302965,0.0006170107,0.0035467779,0.7312581,0.042135485,0.121650524,0.08188514,0.00021131034],"about_ca_topic_score_codex":0.0029669912,"about_ca_topic_score_gemma":0.0045234263,"teacher_disagreement_score":0.013536249,"about_ca_system_score_codex":0.0006579085,"about_ca_system_score_gemma":0.0014468686,"threshold_uncertainty_score":0.045283258},"labels":[],"label_agreement":null},{"id":"W6976965046","doi":"10.60692/kwcnw-nrs71","title":"Integrating Semantics and Neighborhood Information with Graph-Driven Generative Models for Document Retrieval","year":2021,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Semantics (computer science); Generative grammar; Computational linguistics; Joint (building); Natural language; Natural (archaeology); Computational semantics; Natural language generation","score_opus":0.026819208894033702,"score_gpt":0.21204143457446367,"score_spread":0.18522222568042998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976965046","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030943898,0.0027331982,0.9616676,0.00064017286,0.00010110727,0.000101250844,0.0004953208,0.0016615322,0.0016558392],"genre_scores_gemma":[0.75909466,0.0020869158,0.22902003,0.0004248831,0.00037823597,0.0003763448,0.0021665904,0.00078104966,0.0056711864],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921596,0.00037299973,0.000047985985,0.00018012266,0.000114684495,0.00006820301],"domain_scores_gemma":[0.9972154,0.0021061737,0.00014333546,0.00023146071,0.00020602312,0.00009759837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018062047,0.00087264227,0.0019321295,0.0027795646,0.00076894666,0.0015865074,0.00232926,0.0016349279,0.0022487605],"category_scores_gemma":[0.005976652,0.0009700301,0.0020531095,0.0025728755,0.00089725776,0.0039107515,0.0017166603,0.0016422895,0.0012173877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048461807,0.00026204638,0.0038292452,0.00025725664,0.00035942422,0.00018869994,0.00038147075,0.72729236,0.0025119283,0.06984085,0.009568892,0.1850232],"study_design_scores_gemma":[0.000017039383,0.000017099108,0.00013688572,0.000005490362,0.000028027442,0.00002138043,0.000012480958,0.9687205,0.0001534589,0.030386293,0.0004926002,0.000008883718],"about_ca_topic_score_codex":0.016957354,"about_ca_topic_score_gemma":0.029795527,"teacher_disagreement_score":0.016957354,"about_ca_system_score_codex":0.0014250741,"about_ca_system_score_gemma":0.001130361,"threshold_uncertainty_score":0.033717275},"labels":[],"label_agreement":null},{"id":"W6977370478","doi":"10.6084/m9.figshare.19976930","title":"Additional file 1 of CoQUAD: a COVID-19 question answering dataset system, facilitating research, benchmarking, and practice","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto","funders":"","keywords":"Question answering; Exploratory analysis; Exploratory research; Document retrieval","score_opus":0.14672478599036862,"score_gpt":0.37668471526034036,"score_spread":0.22995992926997175,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977370478","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00010726797,0.000008116474,0.00020955273,0.000055054974,0.000012501722,0.000052405478,0.9984622,0.0004104256,0.00068240793],"genre_scores_gemma":[0.0015078811,0.000018942814,0.0018493225,0.00013684986,0.000025198968,0.0008947598,0.99241054,0.00045784944,0.0026986324],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982237,0.0004502155,0.00024794356,0.0005136327,0.00038585256,0.00017850108],"domain_scores_gemma":[0.977749,0.014048974,0.000893379,0.00228942,0.004192611,0.0008266141],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0027132507,0.0010643207,0.0009194648,0.0033592528,0.00095956394,0.0020247195,0.0017712968,0.0013324845,0.59895456],"category_scores_gemma":[0.028870933,0.00051046954,0.0007319312,0.0048354915,0.0003885803,0.0019228377,0.0019299974,0.0012345286,0.22862999],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058796228,0.00003150781,0.00086313253,0.00045395075,0.000009975984,0.000008169883,0.000031688072,0.000085966334,0.00006573935,0.0003396408,0.9950982,0.0029532476],"study_design_scores_gemma":[0.0005640202,0.0000853243,0.0100501655,0.00039087282,0.00003446414,0.000091125694,0.00030232672,0.00090903713,0.00072491285,0.0038176856,0.98296136,0.000068775385],"about_ca_topic_score_codex":0.008031094,"about_ca_topic_score_gemma":0.019888623,"teacher_disagreement_score":0.59895456,"about_ca_system_score_codex":0.0016019737,"about_ca_system_score_gemma":0.0026200837,"threshold_uncertainty_score":0.57204264},"labels":[],"label_agreement":null},{"id":"W6977703902","doi":"10.7275/scil.3141","title":"Similarity, Transformation and the Newly Found Invariance of Influence Functions","year":2025,"lang":"en","type":"article","venue":"University of Massachusetts (UMass) Amherst","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Sentence; Semantics (computer science); Task (project management); Similarity (geometry); Point (geometry); Contrast (vision); Encoding (memory); Transformation (genetics)","score_opus":0.009507513689525567,"score_gpt":0.19586375046241386,"score_spread":0.1863562367728883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977703902","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47139803,0.00047226317,0.51341605,0.00061901176,0.0000640698,0.00007213288,0.00025010598,0.0011461005,0.012562221],"genre_scores_gemma":[0.9771606,0.00010574352,0.021038665,0.00005682738,0.000035923083,0.00003384314,0.00016737574,0.00017015249,0.0012308984],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99858403,0.0005464985,0.00007572687,0.0004122548,0.00026926483,0.00011229221],"domain_scores_gemma":[0.9929616,0.0040057297,0.0006737613,0.0015494411,0.00052856596,0.00028091957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017755674,0.00043344696,0.0004847264,0.0009152536,0.00040675278,0.001606944,0.0005317843,0.0006501448,0.002860571],"category_scores_gemma":[0.021148998,0.00026733725,0.00081011513,0.00058805564,0.0019373021,0.003572048,0.0012056999,0.0011353403,0.00045592478],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009936943,0.00026949338,0.020375488,0.00047718323,0.00029395547,0.00071234343,0.004382721,0.06785717,0.14003131,0.3243764,0.0029129165,0.43731734],"study_design_scores_gemma":[0.000056929504,0.00063833554,0.041561045,0.000032435968,0.0001290661,0.0009617254,0.00056658126,0.41451636,0.047673922,0.48804253,0.0056957253,0.00012538467],"about_ca_topic_score_codex":0.0015719322,"about_ca_topic_score_gemma":0.0009759543,"teacher_disagreement_score":0.002860571,"about_ca_system_score_codex":0.0005615847,"about_ca_system_score_gemma":0.00039363882,"threshold_uncertainty_score":0.009569585},"labels":[],"label_agreement":null},{"id":"W6979293534","doi":"","title":"Syntactic and Semantic Control of Large Language Models via Sequential Monte Carlo","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Advanced Research","keywords":"Probabilistic logic; Python (programming language); Inference; Language model; Variety (cybernetics); Range (aeronautics); Monte Carlo method","score_opus":0.0179119072032372,"score_gpt":0.2600284982508983,"score_spread":0.2421165910476611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979293534","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009463482,0.00008362335,0.98232174,0.00024974704,0.000030996966,0.00008997438,0.00012657988,0.005792139,0.0018416976],"genre_scores_gemma":[0.38609165,0.00015496356,0.606569,0.00036677127,0.00007902139,0.0005955188,0.0006970384,0.0024541605,0.0029919546],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975237,0.0012826137,0.00011570548,0.00044103977,0.0005109739,0.00012592309],"domain_scores_gemma":[0.9907726,0.007284303,0.0002968438,0.00096391793,0.0005260811,0.00015636408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004338199,0.000899365,0.00075621915,0.0006708954,0.00073181366,0.0018693365,0.0025270225,0.0011246043,0.0057045654],"category_scores_gemma":[0.017765714,0.0008121549,0.0013612527,0.0005349447,0.0018158781,0.0024344886,0.0023643293,0.0024647885,0.0016283791],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029297947,0.00015672194,0.0021879938,0.00028289115,0.00012943051,0.00023781747,0.0005301112,0.7559004,0.00998241,0.12386739,0.005999656,0.10043213],"study_design_scores_gemma":[0.000024856025,0.000010401461,0.00004184751,0.0000065820022,0.0000069609227,0.000013084496,0.000010821814,0.96478444,0.0018607188,0.031672794,0.0015615433,0.000005942511],"about_ca_topic_score_codex":0.0053647896,"about_ca_topic_score_gemma":0.0102019785,"teacher_disagreement_score":0.0057045654,"about_ca_system_score_codex":0.001567332,"about_ca_system_score_gemma":0.0020840901,"threshold_uncertainty_score":0.0229429},"labels":[],"label_agreement":null},{"id":"W6979310868","doi":"","title":"The Great Nugget Recall: Automating Fact Extraction and RAG Evaluation with Large Language Models","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute for Information and Communications Technology Promotion; Natural Sciences and Engineering Research Council of Canada; National Institute of Standards and Technology; Ministry of Science and ICT, South Korea","keywords":"Context (archaeology); Focus (optics); Quality (philosophy); Information extraction; Work (physics); Language model; Question answering; Track (disk drive)","score_opus":0.03968283115839261,"score_gpt":0.3117375363597709,"score_spread":0.2720547052013783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979310868","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11515551,0.0014493556,0.82843345,0.0012514392,0.00019072917,0.00092623866,0.0021783835,0.045193758,0.0052211056],"genre_scores_gemma":[0.47013554,0.00021781988,0.5186014,0.0004893755,0.00008065152,0.00062559644,0.005283726,0.0027114484,0.0018544084],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9545275,0.029972225,0.0024195001,0.005291311,0.0068986495,0.000890835],"domain_scores_gemma":[0.90038687,0.06653742,0.0049957624,0.018612588,0.008226438,0.0012409792],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03939525,0.0021375103,0.0014552244,0.0045538303,0.0014515975,0.005792429,0.003264736,0.0024475383,0.0028885605],"category_scores_gemma":[0.09198347,0.0013051641,0.0019038544,0.0020255127,0.0018005094,0.008818558,0.0054249642,0.0036288938,0.0019331661],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012785108,0.0008487383,0.03652295,0.0017775242,0.0012592603,0.00060747884,0.006877974,0.086053245,0.050975177,0.020665098,0.03191794,0.7612161],"study_design_scores_gemma":[0.00020613454,0.00073942845,0.015522783,0.00029100222,0.00038109097,0.0005040161,0.0010740067,0.85940903,0.06736725,0.024501942,0.02968282,0.0003205538],"about_ca_topic_score_codex":0.010192714,"about_ca_topic_score_gemma":0.017417798,"teacher_disagreement_score":0.03939525,"about_ca_system_score_codex":0.0027246003,"about_ca_system_score_gemma":0.0034866752,"threshold_uncertainty_score":0.20834452},"labels":[],"label_agreement":null},{"id":"W6979329680","doi":"","title":"Support Evaluation for the TREC 2024 RAG Track: Comparing Human versus LLM Judges","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute for Information and Communications Technology Promotion; Natural Sciences and Engineering Research Council of Canada; National Institute of Standards and Technology; Ministry of Science and ICT, South Korea","keywords":"Scratch; Human error; Language model; Documentation","score_opus":0.18473266924482978,"score_gpt":0.3754335193551211,"score_spread":0.1907008501102913,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979329680","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94859695,0.0014553961,0.023841871,0.001342592,0.0005542757,0.0008812468,0.00243061,0.0031099434,0.017787082],"genre_scores_gemma":[0.96638805,0.00019415926,0.025482174,0.00052135444,0.00019419778,0.0005576811,0.0034439848,0.00039438118,0.0028239798],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.95530504,0.028955825,0.002961082,0.003944596,0.00778374,0.0010497421],"domain_scores_gemma":[0.81748503,0.13236405,0.008513779,0.012085348,0.025496595,0.0040551554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03884235,0.00094102416,0.0010352421,0.0024792901,0.0019602028,0.0033150984,0.0016615251,0.0021078356,0.00281606],"category_scores_gemma":[0.16044182,0.00036315032,0.0007303057,0.001396848,0.0011960433,0.0026835024,0.002894841,0.0015480651,0.0021171044],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011970133,0.0025559815,0.18919331,0.003865482,0.0012313916,0.001504748,0.07525199,0.016859638,0.055862162,0.004013392,0.106178544,0.5315133],"study_design_scores_gemma":[0.003349658,0.015154894,0.4715006,0.0013906003,0.0011916445,0.0033101693,0.046764985,0.21446411,0.07846992,0.015586566,0.14719117,0.0016256269],"about_ca_topic_score_codex":0.004180469,"about_ca_topic_score_gemma":0.00685893,"teacher_disagreement_score":0.03884235,"about_ca_system_score_codex":0.0012560275,"about_ca_system_score_gemma":0.0013849548,"threshold_uncertainty_score":0.20542043},"labels":[],"label_agreement":null},{"id":"W6979343253","doi":"","title":"Can LLMs Reason Abstractly Over Math Word Problems Without CoT? Disentangling Abstract Formulation From Arithmetic Computation","year":2025,"lang":"en","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Mitacs; Canadian Institute for Advanced Research; Nvidia","keywords":"Conflation; Computation; Causal reasoning; Word (group theory); Automated reasoning; Mental arithmetic; Multiplication (music)","score_opus":0.03054045886540415,"score_gpt":0.27543538240461196,"score_spread":0.2448949235392078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6979343253","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26239887,0.0006995482,0.6920386,0.0034911227,0.00027277687,0.00023055193,0.0021907915,0.016314872,0.02236297],"genre_scores_gemma":[0.7231978,0.00024644297,0.2667153,0.0006232532,0.000050306568,0.00017109005,0.0028328365,0.0019134507,0.0042494657],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99564064,0.0018279243,0.0003030929,0.0007038289,0.001229782,0.00029472043],"domain_scores_gemma":[0.9772613,0.013420397,0.0010101272,0.0056012054,0.0021154818,0.00059146015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005298142,0.0010910019,0.0009119962,0.0010590749,0.0005753233,0.005095216,0.0020545686,0.0014618217,0.01002743],"category_scores_gemma":[0.04829535,0.0006442177,0.0018327753,0.0008879835,0.0022120746,0.013563074,0.004215698,0.0033338638,0.00294879],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014775103,0.0005855209,0.03469262,0.0018662107,0.00036087135,0.00045443818,0.0045317337,0.19394916,0.026577668,0.34430742,0.023031885,0.36816505],"study_design_scores_gemma":[0.00006635393,0.00022188235,0.002373902,0.00015287616,0.00011070432,0.00013334033,0.0008789994,0.6034605,0.018171137,0.35597247,0.018384162,0.00007375427],"about_ca_topic_score_codex":0.0055269683,"about_ca_topic_score_gemma":0.010434401,"teacher_disagreement_score":0.01002743,"about_ca_system_score_codex":0.0018545355,"about_ca_system_score_gemma":0.003126847,"threshold_uncertainty_score":0.033545136},"labels":[],"label_agreement":null},{"id":"W6999866039","doi":"","title":"Distributed prediction of relations for entities: the Easy, the Difficult, and the impossible","year":2017,"lang":"en","type":"article","venue":"Repositori digital de la UPF (Universitat Pompeu Fabra)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Semantics (computer science); Joint (building); Computational semantics; Computational complexity theory","score_opus":0.008070176478722877,"score_gpt":0.21073842281597815,"score_spread":0.20266824633725528,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6999866039","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4179873,0.009619647,0.5418463,0.014670987,0.00074362487,0.00017850929,0.0021803656,0.0018233734,0.010950041],"genre_scores_gemma":[0.9418291,0.0010580975,0.050381735,0.00029215103,0.0007391918,0.00006199331,0.001959404,0.00023964448,0.0034386122],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959214,0.0018318346,0.000209904,0.00129326,0.00044075356,0.00030295004],"domain_scores_gemma":[0.9728358,0.019852366,0.00096117123,0.003922339,0.0011479445,0.0012803858],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0069219763,0.0010106205,0.0017837486,0.0018308295,0.0015934759,0.0044921,0.0025296719,0.0026365558,0.0036843785],"category_scores_gemma":[0.023381358,0.0006287364,0.0010731407,0.0018256276,0.0016839314,0.010612457,0.005141489,0.0040935366,0.0014876951],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032191211,0.0009252948,0.06435301,0.00045691087,0.00049558433,0.00064215146,0.002199051,0.13114595,0.0044934177,0.05375689,0.03241551,0.70589703],"study_design_scores_gemma":[0.00009745128,0.00007529051,0.006370031,0.000053063242,0.00008988882,0.00013843356,0.0005650341,0.84007394,0.001169937,0.14699435,0.00433142,0.00004115441],"about_ca_topic_score_codex":0.008882884,"about_ca_topic_score_gemma":0.013498474,"teacher_disagreement_score":0.008882884,"about_ca_system_score_codex":0.0013265831,"about_ca_system_score_gemma":0.0013915376,"threshold_uncertainty_score":0.036607325},"labels":[],"label_agreement":null},{"id":"W7000541850","doi":"","title":"Focused hierarchical RNNs for conditional sequence processing","year":2018,"lang":"en","type":"article","venue":"Jagiellonian University Repository (Jagiellonian University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"European Regional Development Fund; Canada Research Chairs; Compute Canada; Microsoft Research","keywords":"Security token; Generalization; Sequence (biology); Encoder; Context (archaeology); Recurrent neural network; Embedding; Key (lock)","score_opus":0.023814510767638488,"score_gpt":0.22047415298339149,"score_spread":0.196659642215753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7000541850","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008977733,0.00086344965,0.98196995,0.0002331014,0.00009546539,0.00008592218,0.00060989335,0.004713167,0.0024513335],"genre_scores_gemma":[0.35550454,0.0011969763,0.6212957,0.0005787661,0.00024310141,0.00050936977,0.0046260725,0.0008127795,0.015232701],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99953103,0.00013328661,0.000025946818,0.00017788292,0.00008600748,0.000045758756],"domain_scores_gemma":[0.9991177,0.00044667235,0.000076436714,0.00016181574,0.00016535651,0.000032006672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00090065645,0.001398986,0.0005864624,0.0007503554,0.0002824719,0.0006224331,0.0017145191,0.0010152081,0.007829779],"category_scores_gemma":[0.002842046,0.00055894407,0.0007821736,0.000758652,0.0004316749,0.0016481738,0.00082648796,0.001872265,0.0028632365],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018379942,0.00015260646,0.00079513463,0.00033379745,0.00015526042,0.00012990693,0.00012614256,0.5440724,0.020239344,0.03848831,0.019508472,0.3758148],"study_design_scores_gemma":[0.0000060753882,0.00001641549,0.0001210941,0.0000084537505,0.000008229036,0.000012190129,0.000004197004,0.9875773,0.0021682722,0.008568424,0.0015040784,0.000005308081],"about_ca_topic_score_codex":0.009396349,"about_ca_topic_score_gemma":0.01661568,"teacher_disagreement_score":0.009396349,"about_ca_system_score_codex":0.0012095298,"about_ca_system_score_gemma":0.0010492392,"threshold_uncertainty_score":0.026193261},"labels":[],"label_agreement":null},{"id":"W7005561867","doi":"","title":"Retrieval-augmented text generation with domain-specific large language models fine-tuning","year":2024,"lang":"en","type":"dissertation","venue":"Repository of the University of Ljubljana (University of Ljubljana)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Context (archaeology); Generator (circuit theory); Question answering; Component (thermodynamics); Language model; Architecture; Context model; Labrador Retriever","score_opus":0.013349411071739578,"score_gpt":0.1882553743952043,"score_spread":0.1749059633234647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7005561867","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10337021,0.0003579725,0.8295131,0.00030399032,0.00014508079,0.00054590293,0.00082128204,0.06122522,0.0037172697],"genre_scores_gemma":[0.4850241,0.00013382212,0.5039096,0.0003933172,0.000057644942,0.0007325497,0.0028385296,0.002499157,0.004411268],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99869746,0.00055891974,0.000111173715,0.00036603186,0.00018029106,0.00008614597],"domain_scores_gemma":[0.99713266,0.0015776145,0.00011242657,0.0005719247,0.0005364544,0.000068914196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019729892,0.001322424,0.0007636914,0.0006803861,0.00026721266,0.0010136516,0.0019384328,0.0012496897,0.0038349398],"category_scores_gemma":[0.008020639,0.0006886353,0.0010372323,0.00043663834,0.0005548222,0.0021070493,0.0017429431,0.0016525263,0.0030003337],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006766956,0.0006895124,0.0032330686,0.00068863446,0.00017920863,0.00046749343,0.00080966274,0.40966877,0.10998982,0.004986202,0.011504023,0.45710695],"study_design_scores_gemma":[0.00011536158,0.00014054406,0.00038915526,0.000011153234,0.000041560637,0.00009304903,0.00005794607,0.96145123,0.031837214,0.0021719688,0.0036524127,0.00003844625],"about_ca_topic_score_codex":0.0034264438,"about_ca_topic_score_gemma":0.0043039974,"teacher_disagreement_score":0.0038349398,"about_ca_system_score_codex":0.00068243686,"about_ca_system_score_gemma":0.0009467992,"threshold_uncertainty_score":0.012829125},"labels":[],"label_agreement":null},{"id":"W7008585676","doi":"","title":"Comparing information extraction between&#13;\\ninstance-based data models and relational data&#13;\\nmodels","year":2023,"lang":"en","type":"dissertation","venue":"Memorial University Research Repository (Memorial University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Memorial University of Newfoundland","keywords":"Complement (music); Representation (politics); Information extraction; Data modeling; Data extraction; Data model (GIS); External Data Representation; Key (lock); Data type; Information model","score_opus":0.15516157677851863,"score_gpt":0.31449010470504346,"score_spread":0.15932852792652483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7008585676","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7320619,0.003582131,0.24012555,0.0033676461,0.00017007769,0.0018966058,0.003974774,0.0046399166,0.010181386],"genre_scores_gemma":[0.6638562,0.0018533316,0.3232143,0.0007233855,0.000038029633,0.001137007,0.007392087,0.0002871861,0.0014984866],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9904591,0.0050171614,0.0013683134,0.0012081838,0.0017030952,0.00024412981],"domain_scores_gemma":[0.8537039,0.13092503,0.0038021172,0.007477519,0.003563365,0.0005281111],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014743705,0.0009578107,0.0010185427,0.0036624486,0.00054927915,0.0044141244,0.0017996142,0.0013912308,0.0016916293],"category_scores_gemma":[0.09166316,0.0006907142,0.0020628653,0.0038320103,0.00075323007,0.015279191,0.003267398,0.0014633267,0.0006608388],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006281653,0.0025291652,0.06544095,0.0072266883,0.0017637871,0.0009575298,0.016413447,0.022548227,0.025282979,0.027919045,0.009649152,0.81398743],"study_design_scores_gemma":[0.0014293768,0.004976031,0.09478578,0.0027597584,0.00322556,0.0026963532,0.01999829,0.6206835,0.10948897,0.056587685,0.08256821,0.0008005676],"about_ca_topic_score_codex":0.0033809333,"about_ca_topic_score_gemma":0.003650341,"teacher_disagreement_score":0.014743705,"about_ca_system_score_codex":0.0018863649,"about_ca_system_score_gemma":0.0013192219,"threshold_uncertainty_score":0.07797307},"labels":[],"label_agreement":null},{"id":"W7009646353","doi":"","title":"Explore the In-context Learning Capability of Large Language Models","year":2024,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Blackberry (Canada)","funders":"","keywords":"Benchmark (surveying); Context (archaeology); Embodied cognition; Domain (mathematical analysis); Software deployment; Natural language understanding; Natural language","score_opus":0.017075753417247628,"score_gpt":0.2242240756011204,"score_spread":0.20714832218387277,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7009646353","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2107245,0.0035517656,0.76708424,0.0029749775,0.00018554505,0.00015894261,0.0009865992,0.007610261,0.006723218],"genre_scores_gemma":[0.8118928,0.0006485161,0.18189104,0.0007394475,0.00012903169,0.00012138629,0.002030965,0.00039387922,0.002152976],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99855906,0.0007373216,0.000050025865,0.00039875426,0.00017691482,0.00007797801],"domain_scores_gemma":[0.9939704,0.0045863204,0.00018383838,0.00080885197,0.00030012868,0.00015039866],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024451623,0.0012246176,0.00067847554,0.0008712399,0.0006303747,0.0018452869,0.0019207845,0.0014489423,0.0018876061],"category_scores_gemma":[0.013531964,0.00052611786,0.0010307939,0.00069031876,0.0007738523,0.0053769755,0.0026200586,0.0033497135,0.00083704805],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041230078,0.0004528722,0.008717913,0.00040069045,0.00031801185,0.00024066384,0.0009540678,0.54120904,0.009019199,0.027964307,0.008075449,0.40223542],"study_design_scores_gemma":[0.000013008332,0.00007580912,0.00033778732,0.000016607757,0.000017344864,0.00002872945,0.00008939442,0.9716779,0.0020583102,0.023847992,0.0018257145,0.0000113999295],"about_ca_topic_score_codex":0.0056847213,"about_ca_topic_score_gemma":0.01084902,"teacher_disagreement_score":0.0056847213,"about_ca_system_score_codex":0.0009993493,"about_ca_system_score_gemma":0.0011118344,"threshold_uncertainty_score":0.0129314065},"labels":[],"label_agreement":null},{"id":"W7015239951","doi":"","title":"Simple Convolutional Neural Networks with Linguistically-Annotated Input for Answer Selection in Question Answering","year":2018,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Waterloo","keywords":"Nucleofection; Gestational period; Fusible alloy; Hyporeflexia; TSG101; Dysgeusia; Proteogenomics","score_opus":0.009709399239720461,"score_gpt":0.22043778453560448,"score_spread":0.210728385295884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015239951","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08368118,0.0033743365,0.8734025,0.0022668908,0.00039646338,0.00026597045,0.0034922252,0.014498572,0.018621897],"genre_scores_gemma":[0.75272256,0.00090755796,0.21830152,0.0006934702,0.00024980347,0.0002974535,0.010157999,0.00049051887,0.016179118],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990645,0.00031155517,0.00004802476,0.00033177595,0.00014480801,0.00009930589],"domain_scores_gemma":[0.9985734,0.0007287566,0.00008522096,0.00029957498,0.00025270804,0.00006027795],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013349424,0.0009057766,0.00051074283,0.0009957132,0.00059340213,0.0011989739,0.0017177314,0.0014442544,0.00853335],"category_scores_gemma":[0.0059472397,0.00041508107,0.0006486035,0.0010557093,0.0005766846,0.003684156,0.0013327168,0.0016555821,0.0039018516],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010064393,0.0005054258,0.0049485886,0.00064732233,0.000182835,0.00032183548,0.0007839682,0.102363035,0.032753803,0.03895963,0.0451828,0.7723442],"study_design_scores_gemma":[0.000025226389,0.00005877871,0.0012416994,0.00004355787,0.00004243235,0.000048054375,0.00007755875,0.9433911,0.009566164,0.032235567,0.013248653,0.000021247146],"about_ca_topic_score_codex":0.010243027,"about_ca_topic_score_gemma":0.023479648,"teacher_disagreement_score":0.010243027,"about_ca_system_score_codex":0.0013599569,"about_ca_system_score_gemma":0.0010390348,"threshold_uncertainty_score":0.02854693},"labels":[],"label_agreement":null},{"id":"W7015767590","doi":"","title":"TESA: A task in entity semantic aggregation for abstractive automatic summarization","year":2021,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Task (project management); Automatic summarization; Semantics (computer science); Feature (linguistics); Semantic feature","score_opus":0.0165575354271419,"score_gpt":0.2519691149439467,"score_spread":0.2354115795168048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015767590","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028256778,0.0015340826,0.83218294,0.0018103434,0.0009448581,0.0009373819,0.021520363,0.1046765,0.008136762],"genre_scores_gemma":[0.08391488,0.00058464287,0.83866286,0.0003172588,0.0004576928,0.00083672383,0.056667157,0.005430605,0.01312818],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99774027,0.000702741,0.0002531784,0.00072084373,0.00042527288,0.00015763275],"domain_scores_gemma":[0.9946278,0.002944255,0.0002547823,0.0007358855,0.001150522,0.0002868665],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029551436,0.0028390733,0.0021477318,0.0045788037,0.0028021939,0.0038976339,0.0019008801,0.0022746832,0.02078843],"category_scores_gemma":[0.009735985,0.001058144,0.0024609272,0.003676664,0.0005334413,0.005037834,0.0030776907,0.002665064,0.017083725],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082725694,0.00036563448,0.0017819671,0.0009974204,0.0003149902,0.0003615831,0.0009989073,0.006190612,0.023069937,0.009011175,0.2581562,0.6979243],"study_design_scores_gemma":[0.00050387357,0.00062504667,0.0064777834,0.00025723138,0.00063058274,0.00062062714,0.002504481,0.6075808,0.06617033,0.056564428,0.25784364,0.00022125851],"about_ca_topic_score_codex":0.0053325696,"about_ca_topic_score_gemma":0.009155345,"teacher_disagreement_score":0.02078843,"about_ca_system_score_codex":0.0008866501,"about_ca_system_score_gemma":0.0019491294,"threshold_uncertainty_score":0.069544256},"labels":[],"label_agreement":null},{"id":"W7015919921","doi":"","title":"Understanding and evaluating neural abstractive summarizers using contrastive examples","year":2019,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Vocabulary; Set (abstract data type); Artificial neural network; Benchmark (surveying); Semantics (computer science); Focus (optics)","score_opus":0.14561371146433613,"score_gpt":0.3125817298623217,"score_spread":0.16696801839798556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015919921","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.80660933,0.0063324226,0.1693999,0.00089486886,0.00037447028,0.00073148217,0.0031466014,0.0063749133,0.0061360607],"genre_scores_gemma":[0.83283377,0.0010286864,0.15223502,0.00025385615,0.00016036394,0.00032083696,0.00997819,0.00015422844,0.0030351735],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979911,0.0007157882,0.00028738283,0.00052425155,0.00040557986,0.000075876305],"domain_scores_gemma":[0.98761535,0.009036237,0.0011036231,0.00058772584,0.0013947451,0.00026227074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038718197,0.0012867746,0.00092015666,0.0018992687,0.00040654343,0.0014837782,0.001070023,0.0014829738,0.0018158564],"category_scores_gemma":[0.021322459,0.00025313706,0.00066123094,0.0008399228,0.00041380734,0.0024269498,0.00072562054,0.0011557742,0.00068508874],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027692292,0.0010331366,0.0151179675,0.0024733802,0.0009541489,0.0004704083,0.00086672214,0.27975205,0.03873109,0.0037613413,0.0108607225,0.6432097],"study_design_scores_gemma":[0.00025574618,0.0019319614,0.0075953607,0.00011457939,0.0002473303,0.00016892348,0.0003701983,0.94939494,0.03319502,0.003342115,0.003334794,0.000048928618],"about_ca_topic_score_codex":0.002710968,"about_ca_topic_score_gemma":0.004685025,"teacher_disagreement_score":0.0038718197,"about_ca_system_score_codex":0.0011223085,"about_ca_system_score_gemma":0.0006268573,"threshold_uncertainty_score":0.020476341},"labels":[],"label_agreement":null},{"id":"W7020768153","doi":"","title":"Linguistic change in a nonstandard dialect: phonological studies in the history&#13;\\nof English in Ireland","year":2013,"lang":"en","type":"other","venue":"Edinburgh Research Archive (University of Edinburgh)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Scots; Focus (optics); Interpretation (philosophy); Population; Phonology; Norwegian; Sound change; Welsh; Variation (astronomy)","score_opus":0.13366899685682093,"score_gpt":0.32220472139634754,"score_spread":0.1885357245395266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7020768153","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8636639,0.005395238,0.00054029253,0.0026230635,0.00017842071,0.000023105114,0.00016387738,0.000018414152,0.12739363],"genre_scores_gemma":[0.9905053,0.0012555463,0.0001718331,0.00019826514,0.000050916668,0.00001098133,0.000050279046,0.000023643925,0.007733235],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.998923,0.00033147595,0.00008202979,0.00024341804,0.00016714424,0.00025285754],"domain_scores_gemma":[0.9984434,0.00068557047,0.00022724798,0.00014013433,0.00034597487,0.0001576521],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017075002,0.00014014322,0.00028508733,0.0025638079,0.0029333837,0.0048662466,0.0009024922,0.0005965377,0.0029861876],"category_scores_gemma":[0.003267196,0.00020592356,0.0001910562,0.0029016675,0.006970374,0.0029303632,0.0025910076,0.0012602984,0.00036343568],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013546852,0.00009826262,0.044371452,0.00040676363,0.000023845594,0.0012782698,0.80026966,0.0001608673,0.002798984,0.038353667,0.002975843,0.10912689],"study_design_scores_gemma":[0.000010243275,0.00016224774,0.42932764,0.0003231728,0.000030675936,0.00079675944,0.44300255,0.00019797195,0.0009039711,0.0031934087,0.12198064,0.0000707057],"about_ca_topic_score_codex":0.10406987,"about_ca_topic_score_gemma":0.23366404,"teacher_disagreement_score":0.10406987,"about_ca_system_score_codex":0.007051435,"about_ca_system_score_gemma":0.002468017,"threshold_uncertainty_score":0.20692825},"labels":[],"label_agreement":null},{"id":"W7020803504","doi":"","title":"Natural regeneration of black spruce (Picea mariana (Mill) B.S.P.) on lowland clearcut strips near Shebandowan, Ontario / by Krzysztof Sas-Zmudzinski","year":2017,"lang":"en","type":"other","venue":"Knowledge Commons (Lakehead University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Stocking; Hardwood; Natural regeneration; Competition (biology); Ecological succession; Silviculture; Black spruce; Regeneration (biology); Clearcutting","score_opus":0.021946203613036636,"score_gpt":0.22241808941070926,"score_spread":0.2004718857976726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7020803504","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9993124,0.000051886316,0.000065895096,0.000004420932,7.6317053e-7,0.000005397714,0.00008967067,0.0000042369725,0.00046532313],"genre_scores_gemma":[0.996698,0.000073107,0.0003766468,0.000008758682,8.7464406e-7,0.000010994717,0.00046590954,0.000005125775,0.0023606194],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.99990916,0.000007675766,0.0000037246302,0.000029562609,0.000026437452,0.000023442217],"domain_scores_gemma":[0.99989545,0.000011809154,0.00002022013,0.0000073856677,0.000024094441,0.000040985477],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00012061765,0.0001652007,0.00013304448,0.00034016347,0.0005359741,0.00024356517,0.00027192195,0.000076453805,0.00069085055],"category_scores_gemma":[0.00008579113,0.00016058654,0.00014297751,0.00015775695,0.0002422349,0.00009763083,0.00019612504,0.00013956579,0.00021658928],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006483616,0.00021178489,0.26369372,0.00012815888,0.000060413622,0.0005398043,0.0013877957,0.0006022888,0.70487237,0.00013119883,0.00049289427,0.02723128],"study_design_scores_gemma":[0.0000059975596,0.000119175165,0.995782,0.000004280915,0.000008026609,0.000043234148,0.00018217173,0.00028318315,0.002959254,0.000010053827,0.00059918774,0.0000034427987],"about_ca_topic_score_codex":0.23837213,"about_ca_topic_score_gemma":0.6863717,"teacher_disagreement_score":0.76162785,"about_ca_system_score_codex":0.0016495043,"about_ca_system_score_gemma":0.0007975758,"threshold_uncertainty_score":0.47396934},"labels":[],"label_agreement":null},{"id":"W7021012081","doi":"","title":"\\n L’Histoire des plus grands succès du cinéma Cécile Berger Montréal : Éditions internationales Alain Stanké, 2003 192 pages","year":2004,"lang":"fr","type":"article","venue":"Érudit (Université de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Perspective (graphical); Interpretation (philosophy); Meaning (existential); Relation (database); Feature (linguistics)","score_opus":0.009904548930598622,"score_gpt":0.16771636580780885,"score_spread":0.15781181687721024,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7021012081","genre_codex":"review","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0063268635,0.48009887,0.008324535,0.13657777,0.047928713,0.00008655313,0.0017668419,0.0007974146,0.3180925],"genre_scores_gemma":[0.064742826,0.124220744,0.003918197,0.0066498276,0.01230805,0.00009650556,0.0010220356,0.00068542117,0.78635633],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99925536,0.0001605463,0.000019538016,0.00009150157,0.00034420114,0.00012880663],"domain_scores_gemma":[0.99891865,0.00023726925,0.00005995314,0.00007333126,0.0005377961,0.00017308471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010474036,0.0012423777,0.00048702979,0.003011052,0.004786212,0.0073264674,0.0010370027,0.0014333015,0.032416657],"category_scores_gemma":[0.003187822,0.00051684614,0.00039876334,0.0035183544,0.0028804275,0.0028656323,0.0015888104,0.0024071923,0.005676622],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020374875,0.000009726851,0.00043621418,0.00015740046,0.000008663922,0.000053978903,0.0028787584,0.00012402232,0.00036416395,0.035269957,0.9065512,0.05412565],"study_design_scores_gemma":[0.0000011829425,0.0000019868305,0.0008972651,0.000110039786,0.000003405986,0.000040429368,0.00066124694,0.000037196933,0.00009141849,0.00077701337,0.9973726,0.0000061212],"about_ca_topic_score_codex":0.5921474,"about_ca_topic_score_gemma":0.739239,"teacher_disagreement_score":0.5921474,"about_ca_system_score_codex":0.0137593085,"about_ca_system_score_gemma":0.010566035,"threshold_uncertainty_score":0.820509},"labels":[],"label_agreement":null},{"id":"W7021148780","doi":"","title":"Not an Activist?: Ableism Meets Ageism in the Canadian Media","year":2018,"lang":"en","type":"article","venue":"Project Muse (Johns Hopkins University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ableism; Prejudice (legal term); Racism; Social media; Immigration; Meaning (existential); Narrative; Perspective (graphical)","score_opus":0.03534721765877211,"score_gpt":0.2334121754577838,"score_spread":0.1980649577990117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7021148780","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4881046,0.006738843,0.016210832,0.074626386,0.000770428,0.00015134212,0.0035064437,0.00023980913,0.40965125],"genre_scores_gemma":[0.9855773,0.0009255687,0.0012897499,0.0004988345,0.00010488632,0.000020426627,0.00027258683,0.000064108994,0.01124654],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9976907,0.000675339,0.000046671168,0.0002997672,0.00060495816,0.00068256597],"domain_scores_gemma":[0.99405104,0.0031556906,0.0003512908,0.00026496145,0.0013001635,0.0008769177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039165253,0.00043211094,0.0004022733,0.0035105909,0.012016596,0.01621141,0.0014750978,0.0017012171,0.009494471],"category_scores_gemma":[0.013326151,0.0003265158,0.00040609782,0.0053873416,0.007424597,0.0077427337,0.002850476,0.0024718822,0.0004990479],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002805403,0.0001154342,0.05624426,0.00025586967,0.00005979214,0.00039070443,0.1751297,0.0029519214,0.0007854018,0.6309501,0.05082085,0.082015514],"study_design_scores_gemma":[0.00004604912,0.00004232553,0.074980855,0.00058097823,0.00015567434,0.00025034527,0.3515773,0.01647344,0.0012214488,0.12609383,0.4283561,0.0002217143],"about_ca_topic_score_codex":0.96881145,"about_ca_topic_score_gemma":0.9740563,"teacher_disagreement_score":0.050721195,"about_ca_system_score_codex":0.050721195,"about_ca_system_score_gemma":0.042430844,"threshold_uncertainty_score":0.36800975},"labels":[],"label_agreement":null},{"id":"W7021158969","doi":"","title":"NLL Draft: EP 14 ft. Derek Keenan and Austin Madronic","year":2022,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Period (music); Subject (documents)","score_opus":0.00629075343674189,"score_gpt":0.17873340681145936,"score_spread":0.17244265337471745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7021158969","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00047850152,0.00113646,0.0042951587,0.032568198,0.022884557,0.00039161998,0.04339324,0.0049208766,0.8899314],"genre_scores_gemma":[0.002353764,0.0004448547,0.0011315885,0.003677607,0.0014071309,0.00018034416,0.016045399,0.0021877824,0.9725714],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99860543,0.00025749268,0.000075762015,0.00016652096,0.0007666198,0.00012817222],"domain_scores_gemma":[0.99322695,0.0012739907,0.000117534255,0.0005741515,0.0044023427,0.00040510492],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0028596881,0.0006935988,0.0007373871,0.0015168099,0.0023680774,0.007967896,0.00092279975,0.0020065096,0.66545814],"category_scores_gemma":[0.020041564,0.0005366274,0.0005902014,0.0017717758,0.0007269671,0.0033008314,0.002404615,0.003021235,0.65864533],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000005322163,0.0000022706445,0.000012539671,0.000014122818,4.897498e-7,0.000004475611,0.000008489243,0.00000975428,0.000018772696,0.00071216276,0.9965444,0.0026671896],"study_design_scores_gemma":[0.000007841659,0.000002992838,0.00015728957,0.000062636034,0.000001233223,0.000008174176,0.0000373092,0.000053141433,0.00006175453,0.00090033637,0.99870336,0.000003971193],"about_ca_topic_score_codex":0.03125769,"about_ca_topic_score_gemma":0.048375905,"teacher_disagreement_score":0.33454186,"about_ca_system_score_codex":0.0038021023,"about_ca_system_score_gemma":0.0034269437,"threshold_uncertainty_score":0.47718334},"labels":[],"label_agreement":null},{"id":"W7021175993","doi":"","title":"Nino Haratischwili i jej literacki obraz Gruzji przedstawiony w utworach prozatorskich","year":2024,"lang":"pl","type":"other","venue":"UMCS Library (Maria Curie-Skłodowska University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"WiLAN (Canada)","funders":"","keywords":"Term (time); Perspective (graphical); Subject (documents); Context (archaeology)","score_opus":0.011631583067283499,"score_gpt":0.19276030962845178,"score_spread":0.1811287265611683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7021175993","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3042114,0.03215312,0.0043917536,0.02556562,0.0036872698,0.0000981538,0.0007181552,0.00041339037,0.62876105],"genre_scores_gemma":[0.77551734,0.007306727,0.0017025474,0.0016388282,0.00042428199,0.00007544838,0.00026109224,0.00025083995,0.21282291],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997476,0.00006216945,0.000010452746,0.000049028542,0.00005472242,0.000075987096],"domain_scores_gemma":[0.9998617,0.000032383326,0.000025396206,0.000011177821,0.000024523699,0.00004478934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00030962567,0.00043792784,0.00025549094,0.0009230122,0.0038327074,0.003907098,0.00026116715,0.00048404376,0.0059163165],"category_scores_gemma":[0.00041041986,0.00015128855,0.00010764335,0.0007944942,0.003804841,0.0013415975,0.0021367192,0.001364566,0.0023270526],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033580547,0.00007220852,0.0065061557,0.0008835424,0.000027518809,0.003039046,0.28361896,0.0003395578,0.011547137,0.37630266,0.1732994,0.14402798],"study_design_scores_gemma":[0.0000042414704,0.000016602527,0.0050320914,0.00014557729,0.0000068231834,0.00064024073,0.019581823,0.000061652114,0.0008929067,0.0031747369,0.9704293,0.000013903558],"about_ca_topic_score_codex":0.00623553,"about_ca_topic_score_gemma":0.020106262,"teacher_disagreement_score":0.00623553,"about_ca_system_score_codex":0.0025459935,"about_ca_system_score_gemma":0.0014939841,"threshold_uncertainty_score":0.01979208},"labels":[],"label_agreement":null},{"id":"W7021273625","doi":"","title":"NEVER GIVE UP {{{{NR PHONE NUMBER++NR TECH SUPPORT PHONE NUMBER}}}&amp;lt;&amp;lt;18&amp;gt;&amp;gt;?844&amp;gt;0?9087^^norton pro tech support phone number Public 0 Share","year":2016,"lang":"en","type":"other","venue":"OSF Preprints (OSF Preprints)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Phone; Service (business); Telephone number; Phone call; Helpline; Mobile phone","score_opus":0.037216111762918024,"score_gpt":0.29104872944999427,"score_spread":0.25383261768707627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7021273625","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03102652,0.0042072646,0.0018471958,0.018629158,0.0075236782,0.00059172785,0.0072030034,0.003701817,0.9252696],"genre_scores_gemma":[0.03087437,0.0023337724,0.000702628,0.008027471,0.0006229179,0.0002735508,0.0019175432,0.00044557097,0.95480216],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99902403,0.00013773216,0.000043030464,0.00010879738,0.0003750663,0.00031144067],"domain_scores_gemma":[0.9932372,0.00053994614,0.00045159837,0.0005768102,0.0010567736,0.0041376967],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0010158274,0.000595711,0.0008629305,0.0005295191,0.0020224066,0.0025432308,0.0010473357,0.0013544132,0.5902495],"category_scores_gemma":[0.008861223,0.00038266272,0.00062061154,0.00038240617,0.00035823695,0.0021495055,0.00251149,0.0016056112,0.49776733],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019765335,0.00032452325,0.006298522,0.0002222605,0.000017045813,0.00020931679,0.00026177798,0.000039898547,0.00043833672,0.00067807577,0.82364106,0.16767152],"study_design_scores_gemma":[0.000053065927,0.00030834158,0.016501883,0.0005321997,0.000032336135,0.0010010982,0.0016406841,0.00018512506,0.0004995288,0.0010789648,0.97811925,0.00004751431],"about_ca_topic_score_codex":0.0017286292,"about_ca_topic_score_gemma":0.0045147017,"teacher_disagreement_score":0.40975052,"about_ca_system_score_codex":0.00046860383,"about_ca_system_score_gemma":0.0011056319,"threshold_uncertainty_score":0.58445936},"labels":[],"label_agreement":null},{"id":"W7021988903","doi":"","title":"Seasons of my life","year":2010,"lang":"en","type":"article","venue":"cIRcle (University of British Columbia)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Reflexive pronoun; Poetry; Ballet; Creative writing; Creative work; Style (visual arts); Curriculum; Theme (computing); Verb","score_opus":0.009018953822434431,"score_gpt":0.17055591992483388,"score_spread":0.16153696610239945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7021988903","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.066103615,0.019937776,0.006350772,0.17549248,0.04382882,0.0002994388,0.0022057001,0.000931673,0.68484974],"genre_scores_gemma":[0.35821968,0.0101253865,0.0041833394,0.044103276,0.005990763,0.0003969322,0.0012297564,0.0010875907,0.5746632],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99724764,0.0011015882,0.00007678196,0.00039401147,0.00073705835,0.00044287226],"domain_scores_gemma":[0.9965178,0.0002098195,0.00026224687,0.0002825539,0.00064447016,0.002083125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016216811,0.00076992845,0.00074517744,0.00072228553,0.009396083,0.008472267,0.00086979574,0.0015684628,0.038468555],"category_scores_gemma":[0.006321142,0.0002850558,0.0005641836,0.0007019369,0.0059181373,0.0059927814,0.00789895,0.0056267236,0.01894104],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009462107,0.00007892086,0.0026102008,0.00026521753,0.000034646862,0.0007321306,0.110939905,0.00006146105,0.0013294972,0.06839553,0.7586852,0.056772754],"study_design_scores_gemma":[0.000003222788,0.000030389627,0.0012358623,0.00015949059,0.0000035151281,0.00033754276,0.02911884,0.000011652929,0.00009092502,0.002627143,0.9663647,0.000016720314],"about_ca_topic_score_codex":0.005088308,"about_ca_topic_score_gemma":0.008435646,"teacher_disagreement_score":0.038468555,"about_ca_system_score_codex":0.0026619113,"about_ca_system_score_gemma":0.0027044616,"threshold_uncertainty_score":0.12869018},"labels":[],"label_agreement":null},{"id":"W7022133977","doi":"","title":"Sculpture","year":2021,"lang":"en","type":"article","venue":"UND Scholarly Commons (University of North Dakota)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Liquation; Paraphernalia; Subpoena; Work (physics); TSG101; Circumstantial evidence","score_opus":0.02289001890041631,"score_gpt":0.20954240550419978,"score_spread":0.18665238660378347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7022133977","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021391695,0.0023708963,0.0008961811,0.0018807617,0.001961281,0.00007165108,0.0016522877,0.0008564739,0.9881714],"genre_scores_gemma":[0.008241384,0.001364872,0.00076485565,0.0006987841,0.00037609905,0.000024248264,0.0014906022,0.00041051942,0.98662853],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993345,0.000044284967,0.000027360433,0.00008723842,0.00043366276,0.00007289445],"domain_scores_gemma":[0.9992436,0.00006016926,0.00002922705,0.00012485635,0.0003496621,0.0001925271],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00031535752,0.00079064805,0.0004107726,0.0019784048,0.0025511915,0.003645753,0.0010446188,0.00086304476,0.5094171],"category_scores_gemma":[0.0014582826,0.000280467,0.00047453257,0.0015541717,0.00056860514,0.0016016687,0.0028633266,0.0013939104,0.27101278],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038092167,0.000032877484,0.00043441905,0.000161469,0.000005230403,0.00036117499,0.00042962743,0.000098385615,0.0013178408,0.0046480238,0.8406883,0.15178455],"study_design_scores_gemma":[0.0000015433563,0.000006913981,0.00038063963,0.00004324487,9.774282e-7,0.00022937595,0.000095404976,0.000025851004,0.00012421624,0.00024083666,0.9988475,0.000003521664],"about_ca_topic_score_codex":0.004747065,"about_ca_topic_score_gemma":0.015675364,"teacher_disagreement_score":0.5094171,"about_ca_system_score_codex":0.00096960965,"about_ca_system_score_gemma":0.0010866895,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W7023699541","doi":"","title":"O wheel!, or, Thanksgiving thoughts; a sermon, preached in St.Andrew' s church, Montreal, on Wednesday, 18th October, 1865.","year":2014,"lang":"en","type":"article","venue":"QSpace (Queen's University Library)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Subject (documents); Identification (biology); Period (music); Identity (music)","score_opus":0.010894673460515002,"score_gpt":0.19342997488858404,"score_spread":0.18253530142806904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7023699541","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0028677683,0.097403675,0.0034911986,0.3783722,0.06620417,0.00013272115,0.0037828777,0.0016370841,0.4461083],"genre_scores_gemma":[0.0062852073,0.0052485927,0.00045897727,0.0075562587,0.0018410287,0.000014534626,0.00020236305,0.0003011614,0.9780919],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995757,0.00009801708,0.000014762958,0.00008300542,0.00013993915,0.00008860031],"domain_scores_gemma":[0.9990709,0.00014452777,0.000030334242,0.00003453829,0.00030950812,0.00041017128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001195589,0.0008339241,0.00056583923,0.0007602065,0.009527929,0.0045438036,0.0006001516,0.0016700155,0.15594429],"category_scores_gemma":[0.0034578948,0.00049806567,0.0003270778,0.0009989709,0.0023122237,0.0027014678,0.0017692181,0.0032710463,0.054528337],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000122729325,0.000003467006,0.00009194339,0.000014717914,0.0000012285433,0.000028224098,0.00034281964,0.000011432706,0.000056201236,0.0016650988,0.9919973,0.005775413],"study_design_scores_gemma":[0.0000018228668,0.0000036253277,0.00033618743,0.000029070286,9.287094e-7,0.000028165705,0.0006543729,0.00001305801,0.000030427827,0.00028194612,0.9986155,0.000004837614],"about_ca_topic_score_codex":0.18862768,"about_ca_topic_score_gemma":0.5359313,"teacher_disagreement_score":0.81137234,"about_ca_system_score_codex":0.0053103054,"about_ca_system_score_gemma":0.004061885,"threshold_uncertainty_score":0.5216856},"labels":[],"label_agreement":null},{"id":"W7023812322","doi":"","title":"Performance of a Family of Surface Piercing Propellers","year":2001,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Propulsor; Thrust; Torque; Surface (topology); Propeller; Scaling","score_opus":0.032684444845665114,"score_gpt":0.23589647257808188,"score_spread":0.20321202773241676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7023812322","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98422885,0.00022960534,0.012875887,0.000034328834,0.000021795071,0.000046298184,0.00014308633,0.00039979702,0.0020203346],"genre_scores_gemma":[0.9904767,0.00018978224,0.007568053,0.000011792924,0.0000067018705,0.000021776274,0.00041242078,0.000036597652,0.0012760195],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9996786,0.000023581624,0.000015526653,0.000047914848,0.00016000353,0.00007441956],"domain_scores_gemma":[0.99908197,0.00024061951,0.000108650645,0.000120352386,0.00031712535,0.00013128399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00043408666,0.0005777227,0.0005806689,0.0007484444,0.0003941354,0.00049530127,0.0005196074,0.0005584704,0.0019352593],"category_scores_gemma":[0.0011454943,0.00019166873,0.00039123252,0.0006847375,0.00036676155,0.0004157063,0.0003966493,0.00037481348,0.0005692241],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026726683,0.0004578213,0.014674691,0.00049404823,0.00010252153,0.0007939067,0.0003829542,0.08595299,0.67950344,0.0007547128,0.0013699515,0.21284021],"study_design_scores_gemma":[0.00014704546,0.014522994,0.04637239,0.000045952977,0.0000750603,0.0018124967,0.0004583933,0.34235868,0.5870749,0.00055506977,0.006449357,0.00012772267],"about_ca_topic_score_codex":0.00065627776,"about_ca_topic_score_gemma":0.00036937374,"teacher_disagreement_score":0.0019352593,"about_ca_system_score_codex":0.00016967377,"about_ca_system_score_gemma":0.00016843932,"threshold_uncertainty_score":0.0064741373},"labels":[],"label_agreement":null},{"id":"W7023896150","doi":"","title":"Political and social influences on religious school : a historical perspective on Indonesian Islamic school curricula","year":2006,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"McGill University","keywords":"Islam; Indonesian; Curriculum; Politics; Perspective (graphical); Modernization theory; Religious education; Reciprocal","score_opus":0.013092343724295703,"score_gpt":0.259505813105654,"score_spread":0.2464134693813583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7023896150","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.76879597,0.014230987,0.0010216843,0.006817519,0.00019815963,0.00002529069,0.00014873466,0.000018296167,0.20874341],"genre_scores_gemma":[0.99242944,0.005173033,0.00017180033,0.000119455726,0.00006821406,0.000006445748,0.000023797002,0.0000096065,0.0019982057],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.999433,0.00030086946,0.000021564381,0.000062719344,0.00008346234,0.000098285294],"domain_scores_gemma":[0.998643,0.0007063352,0.00029358687,0.000050448285,0.00011551748,0.00019122199],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013085459,0.00014481197,0.00014873505,0.0018574346,0.0035756768,0.0040926267,0.00037951532,0.0004983812,0.0027443445],"category_scores_gemma":[0.0017641766,0.00020671773,0.000104821396,0.0024977552,0.005181069,0.0029235848,0.0014682395,0.0019337235,0.00021607775],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007999525,0.0003042422,0.10048631,0.00037051633,0.00002221862,0.0015097621,0.5712153,0.00055877696,0.0009651242,0.21896996,0.004643052,0.10087478],"study_design_scores_gemma":[0.0000064266505,0.000085557534,0.37581134,0.0013591786,0.00003706166,0.0009999489,0.345958,0.001305705,0.001307992,0.011354585,0.26171702,0.000057217596],"about_ca_topic_score_codex":0.013142981,"about_ca_topic_score_gemma":0.036813103,"teacher_disagreement_score":0.013142981,"about_ca_system_score_codex":0.0042275907,"about_ca_system_score_gemma":0.0014832338,"threshold_uncertainty_score":0.030673504},"labels":[],"label_agreement":null},{"id":"W7023950838","doi":"","title":"QUEST 2.0 : Apuvälinetyytyväisyyttä arvioivan mittarin käyttöönotto ja soveltuvuus Suomessa","year":2012,"lang":"fi","type":"other","venue":"STM:n Hallinnonalan avoin julkaisuarkisto (Julkari)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Government (linguistics); Work (physics); Perspective (graphical); Context (archaeology); Field (mathematics); Identification (biology)","score_opus":0.02600906274724145,"score_gpt":0.26103130208961556,"score_spread":0.2350222393423741,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7023950838","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34305435,0.0014261339,0.022638703,0.005726791,0.0008205053,0.008777641,0.13143472,0.009209193,0.47691193],"genre_scores_gemma":[0.46860218,0.0012927868,0.07533264,0.0022324384,0.00017704922,0.01190323,0.09487203,0.0034149173,0.34217277],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9955597,0.00095258636,0.00034429692,0.00042493895,0.0022299304,0.0004885901],"domain_scores_gemma":[0.9797452,0.0043605105,0.00065921753,0.0009743054,0.011973131,0.0022877082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008575041,0.0008265502,0.0008038208,0.0018857655,0.0014479609,0.0050446433,0.0010747198,0.0008493715,0.054496262],"category_scores_gemma":[0.014085774,0.00045210656,0.0009620746,0.0014543316,0.00094316446,0.0015676193,0.0019015416,0.0014230798,0.019600423],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022932948,0.0013702202,0.13380107,0.0024839872,0.00016221257,0.0004113004,0.016229657,0.00087466073,0.011963532,0.009957726,0.47032455,0.3501278],"study_design_scores_gemma":[0.00034340558,0.00086729584,0.39553574,0.0009982828,0.00011534999,0.00028074524,0.0083598625,0.0024452112,0.0063589597,0.0030261627,0.58136696,0.0003019751],"about_ca_topic_score_codex":0.121364534,"about_ca_topic_score_gemma":0.2957312,"teacher_disagreement_score":0.121364534,"about_ca_system_score_codex":0.0042957724,"about_ca_system_score_gemma":0.009391685,"threshold_uncertainty_score":0.2413162},"labels":[],"label_agreement":null},{"id":"W7024067126","doi":"","title":"On the problem of Exupérian heroism in Merleau-Ponty's phenomenology of perception","year":2006,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Phenomenology (philosophy); Perception; Relation (database)","score_opus":0.017533581652272416,"score_gpt":0.2319959040031209,"score_spread":0.21446232235084847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024067126","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41788602,0.0054003582,0.07724666,0.047602873,0.00026685925,0.00008890079,0.00009942478,0.0001802276,0.4512287],"genre_scores_gemma":[0.9939709,0.00038459428,0.0012226836,0.00024454834,0.00002672925,0.00001885724,0.000009372315,0.000029622668,0.004092674],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99658185,0.0026381717,0.00004661459,0.00026610913,0.0002879511,0.0001793604],"domain_scores_gemma":[0.9955225,0.0035470284,0.00021004317,0.00029963884,0.00022902068,0.00019178857],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004864762,0.00035709405,0.0003497199,0.0015478721,0.005126248,0.0071562775,0.00089434336,0.0021626325,0.0030713286],"category_scores_gemma":[0.009942725,0.00033098724,0.00037349603,0.001932294,0.036348693,0.012711739,0.003613928,0.003144166,0.00024457395],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000023498602,0.000010729425,0.0006060271,0.00003743152,0.0000045030174,0.00017084212,0.41455495,0.00025009934,0.00015349397,0.5756574,0.001787214,0.006743705],"study_design_scores_gemma":[0.000021312553,0.000033804252,0.0030093954,0.00023044959,0.000012224061,0.0006376023,0.304674,0.0023012587,0.00030653953,0.60610735,0.082615785,0.000050224866],"about_ca_topic_score_codex":0.012017917,"about_ca_topic_score_gemma":0.00950511,"teacher_disagreement_score":0.012017917,"about_ca_system_score_codex":0.004964638,"about_ca_system_score_gemma":0.0018047094,"threshold_uncertainty_score":0.036021113},"labels":[],"label_agreement":null},{"id":"W7024166041","doi":"","title":"Public CoLab 2023. Embracing Belfast's riverfront as public space","year":2023,"lang":"en","type":"other","venue":"Research Portal (Queen's University Belfast)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Queen's University","keywords":"Public space; Space (punctuation); Government (linguistics); Key (lock)","score_opus":0.04917069813045644,"score_gpt":0.28276107643266785,"score_spread":0.2335903783022114,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024166041","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010593537,0.00092630915,0.0029606223,0.015592073,0.0023003977,0.00020339337,0.029917829,0.004371481,0.9426686],"genre_scores_gemma":[0.0039609913,0.00029914433,0.0009875224,0.00096213224,0.00011323192,0.000081716644,0.0049017593,0.0016346275,0.9870588],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99907136,0.00014893587,0.000021846528,0.00010611026,0.00036595212,0.0002857427],"domain_scores_gemma":[0.99610525,0.00059450604,0.000115702154,0.00039806712,0.0008385978,0.0019479526],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0019825425,0.0006428993,0.00040685807,0.0009925219,0.003498413,0.010543846,0.001739785,0.004227645,0.7120311],"category_scores_gemma":[0.005113278,0.00070322445,0.0006888995,0.002656568,0.0013677836,0.005999357,0.0047727367,0.0020063634,0.42530233],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026833335,0.0000082475435,0.000079067126,0.00005611031,8.6760014e-7,0.000020966572,0.00014399023,0.000037542108,0.00006543303,0.0046711,0.982866,0.01202378],"study_design_scores_gemma":[0.000003815545,0.0000037620264,0.00019118935,0.000019906187,4.866458e-7,0.0000044642543,0.0001797895,0.00003119848,0.000024828332,0.00037244897,0.99916506,0.0000030473634],"about_ca_topic_score_codex":0.2722068,"about_ca_topic_score_gemma":0.5096482,"teacher_disagreement_score":0.7120311,"about_ca_system_score_codex":0.00892325,"about_ca_system_score_gemma":0.0087607745,"threshold_uncertainty_score":0.54124475},"labels":[],"label_agreement":null},{"id":"W7024198883","doi":"","title":"Remediation of Heavy Metals with Poplar Trees","year":2023,"lang":"en","type":"dissertation","venue":"UVic’s Research and Learning Repository (University of Victoria)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Stormwater; Environmental remediation; Wastewater; Effluent; Heavy metals; Sewage treatment; Population","score_opus":0.029175951488105797,"score_gpt":0.2730867319718598,"score_spread":0.24391078048375398,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024198883","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9925189,0.00039575133,0.004879905,0.00008303202,0.000017048857,0.000055206907,0.00011536202,0.00007463926,0.0018601627],"genre_scores_gemma":[0.982736,0.00090938364,0.009814765,0.000111504996,0.000009478547,0.000051226474,0.00021363914,0.000017009901,0.0061369846],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99978656,0.0000263856,0.000008927218,0.000046507997,0.0000876366,0.000044015967],"domain_scores_gemma":[0.9998771,0.000018178644,0.000030717285,0.0000087764865,0.00004625216,0.000018967046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00022450389,0.00031253655,0.00028649412,0.00047691326,0.0004092414,0.00084074435,0.0003975368,0.00030121338,0.0009699026],"category_scores_gemma":[0.00024032529,0.00012415442,0.000481283,0.00048163248,0.00016468765,0.00030064845,0.0003278874,0.00047489707,0.0004449216],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013149,0.00025783564,0.005946246,0.00025021404,0.000019732128,0.00016260837,0.00017055598,0.0013197319,0.96749073,0.00019594103,0.0001461968,0.023908736],"study_design_scores_gemma":[0.000022733464,0.0041586976,0.033626124,0.000051249513,0.00008389204,0.0003170107,0.0005896258,0.004096967,0.94696593,0.00029071898,0.009777946,0.000019103009],"about_ca_topic_score_codex":0.0027980877,"about_ca_topic_score_gemma":0.0061136764,"teacher_disagreement_score":0.0027980877,"about_ca_system_score_codex":0.0003114069,"about_ca_system_score_gemma":0.00045139305,"threshold_uncertainty_score":0.0055636168},"labels":[],"label_agreement":null},{"id":"W7024249957","doi":"","title":"Project Khepri: Asteroid Mining Project. Final Policy Report","year":2022,"lang":"en","type":"article","venue":"Scholarship@Western (Western University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Space law; Space (punctuation); Outer space; Space industry; Work (physics); Resource (disambiguation); International law","score_opus":0.18221197415803375,"score_gpt":0.34964968592165735,"score_spread":0.1674377117636236,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024249957","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031278008,0.003317346,0.00412567,0.053380445,0.006118128,0.0032390663,0.10533882,0.0026518642,0.81870085],"genre_scores_gemma":[0.007498091,0.0012454843,0.003148132,0.006565693,0.00032541537,0.0019085301,0.036684755,0.00064078294,0.94198316],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99675876,0.0005309193,0.00011771924,0.00038794076,0.0016056459,0.0005990312],"domain_scores_gemma":[0.99744654,0.00028326828,0.00015769099,0.00018309141,0.0013108872,0.00061851135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065707825,0.0008420752,0.00041251234,0.0013922275,0.002157157,0.004783898,0.0023794407,0.0042260797,0.27186748],"category_scores_gemma":[0.004472011,0.0005136634,0.00046611123,0.0016597837,0.0006546751,0.0033380585,0.0041174153,0.0028580965,0.13350728],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011992418,0.00007890972,0.0004106855,0.00020742275,0.0000040450486,0.000094928015,0.00010802304,0.00011391742,0.0003267048,0.008933835,0.9677439,0.021857714],"study_design_scores_gemma":[0.00001949992,0.00001944584,0.0007836039,0.000057137007,0.0000015300327,0.000020024523,0.00011743044,0.000059404534,0.00019453715,0.00060854555,0.99811375,0.0000050469534],"about_ca_topic_score_codex":0.029115569,"about_ca_topic_score_gemma":0.025506161,"teacher_disagreement_score":0.27186748,"about_ca_system_score_codex":0.004348722,"about_ca_system_score_gemma":0.022249518,"threshold_uncertainty_score":0.90948737},"labels":[],"label_agreement":null},{"id":"W7024411014","doi":"","title":"Renewable energy will power up to 30 per cent of Alberta's electricity grid by 2030","year":2015,"lang":"en","type":"other","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Renewable energy; Electricity; Power grid; Power (physics); Electricity generation; Intermittent energy source; Energy (signal processing)","score_opus":0.009312441949670307,"score_gpt":0.2180904042690242,"score_spread":0.2087779623193539,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024411014","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017802272,0.0033111826,0.025989264,0.017846588,0.0030506547,0.00017880906,0.043251816,0.007463157,0.8811063],"genre_scores_gemma":[0.042931475,0.0031087918,0.00938133,0.0007924005,0.00026256745,0.00004549,0.015775194,0.0005539584,0.9271487],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995993,0.000030302046,0.0000066206553,0.0000407565,0.00022342699,0.00009958273],"domain_scores_gemma":[0.9993319,0.000052384436,0.000021730402,0.00005099051,0.00037547434,0.00016752136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007063571,0.0009591002,0.0002876982,0.0014077217,0.0016733418,0.002866908,0.0010824812,0.0010208397,0.08909266],"category_scores_gemma":[0.001423401,0.00029981838,0.00068562693,0.0022658815,0.000390398,0.0012683651,0.0012774288,0.0009887828,0.029331516],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007238626,0.000045638993,0.0032204427,0.00012622318,0.000014302503,0.00004597054,0.0001337284,0.0054617804,0.00080498867,0.021113705,0.8038208,0.16514],"study_design_scores_gemma":[0.00001585404,0.000011299005,0.004298024,0.00007295322,0.00001790761,0.000023237222,0.00045641136,0.014244146,0.00079581945,0.010180215,0.9698646,0.000019490708],"about_ca_topic_score_codex":0.7985138,"about_ca_topic_score_gemma":0.9298612,"teacher_disagreement_score":0.20148617,"about_ca_system_score_codex":0.008977645,"about_ca_system_score_gemma":0.017700944,"threshold_uncertainty_score":0.4053455},"labels":[],"label_agreement":null},{"id":"W7024539609","doi":"","title":"The roles of MUSE1 and MUSE15 in plant innate immunity","year":2016,"lang":"en","type":"other","venue":"cIRcle (University of British Columbia)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Arizona Biomedical Research Commission","keywords":"TSG101; Ubiquitin-Protein Ligases; Identification (biology); Mechanism (biology); Mutation","score_opus":0.007841298324133137,"score_gpt":0.17029190680714898,"score_spread":0.16245060848301585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024539609","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9926682,0.003894708,0.0016800262,0.00006543127,0.000012107307,0.000010369597,0.0002013208,0.000043824533,0.001424174],"genre_scores_gemma":[0.9939155,0.0011765365,0.0019882838,0.000049727692,0.000008288613,0.000009581109,0.00063185865,0.000014081103,0.0022061816],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9998882,0.000009835467,0.000010390047,0.000033598197,0.0000299846,0.000027999835],"domain_scores_gemma":[0.9998783,0.0000110608935,0.00004114169,0.000013753943,0.00001756648,0.000038180155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00017594779,0.0004638373,0.00029459572,0.00023980526,0.00017670993,0.0005299026,0.00027873696,0.0003271059,0.00063166296],"category_scores_gemma":[0.00010915422,0.00013060431,0.0003752019,0.00013203216,0.00017524122,0.00032961275,0.00039786,0.00030491024,0.00026202813],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010559923,0.0000083423965,0.002606906,0.000042624382,0.000006362232,0.00013205242,0.000019696583,0.00009481456,0.99400663,0.000308717,0.000029656183,0.0026385286],"study_design_scores_gemma":[0.000020703823,0.00032832095,0.09381585,0.000026264219,0.000050441948,0.0015369444,0.00016071198,0.0021040451,0.8939444,0.00053509435,0.007451813,0.000025511932],"about_ca_topic_score_codex":0.00022664301,"about_ca_topic_score_gemma":0.00040361407,"teacher_disagreement_score":0.00063166296,"about_ca_system_score_codex":0.00035651645,"about_ca_system_score_gemma":0.00022717935,"threshold_uncertainty_score":0.0025866628},"labels":[],"label_agreement":null},{"id":"W7024716328","doi":"","title":"Search for WZ resonances in the fully leptonic channel using pp collisions at root s=8 TeV with the ATLAS detector","year":2014,"lang":"en","type":"article","venue":"Lancaster EPrints (Lancaster University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fundação para a Ciência e a Tecnologia; Institut National de Physique Nucléaire et de Physique des Particules; Agencia Nacional de Promoción Científica y Tecnológica; Science and Technology Facilities Council; Natural Sciences and Engineering Research Council of Canada; H. Lundbeck A/S; State Atomic Energy Corporation ROSATOM; Centre National pour la Recherche Scientifique et Technique; Georgian National Science Foundation; Centre National de la Recherche Scientifique; Max-Planck-Gesellschaft; Israel Science Foundation; Lundbeckfonden; Leverhulme Trust; General Secretariat for Research and Technology; Ministry of Education, Culture, Sports, Science and Technology; Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Austrian Science Fund; Bundesministerium für Bildung und Forschung; Israeli Centers for Research Excellence; Joint Institute for Nuclear Research; National Science Council; Japan Society for the Promotion of Science; Conselho Nacional de Desenvolvimento Científico e Tecnológico; U.S. Department of Energy; National Natural Science Foundation of China; Fundação de Amparo à Pesquisa do Estado de São Paulo; Bundesministerium für Wissenschaft und Forschung; Javna Agencija za Raziskovalno Dejavnost RS; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; Ministerstwo Edukacji i Nauki; CERN; Deutsche Forschungsgemeinschaft; Services Fédéraux des Affaires Scientifiques, Techniques et Culturelles; Department of Science and Technology, Ministry of Science and Technology, India; European Commission; Comisión Nacional de Investigación Científica y Tecnológica; Danmarks Grundforskningsfond; TRIUMF; Alexander von Humboldt-Stiftung; Türkiye Atom Enerjisi Kurumu; National Science Foundation","keywords":"Large Hadron Collider; Atlas (anatomy); Detector; Atlas detector; Standard Model (mathematical formulation); Channel (broadcasting); Limit (mathematics); Collision","score_opus":0.03913579680085898,"score_gpt":0.2354994904192568,"score_spread":0.19636369361839784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024716328","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9857636,0.0000977859,0.010047779,0.000054470805,0.0000030437434,0.000015883506,0.00054886536,0.0003400065,0.00312847],"genre_scores_gemma":[0.9933119,0.000067477675,0.00487443,0.000016884034,0.000005990681,0.000012526728,0.00086947257,0.000025206633,0.0008160456],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99956995,0.00012883391,0.000017520482,0.000105004656,0.00010765461,0.00007106136],"domain_scores_gemma":[0.9995395,0.00023236837,0.00008014981,0.000055962446,0.000034480552,0.00005763331],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007497082,0.0005979033,0.0006167937,0.0012639042,0.0008657649,0.001301742,0.0007664812,0.0006032513,0.0022525468],"category_scores_gemma":[0.00067068025,0.0005040661,0.00045057834,0.0021096426,0.00027902625,0.0007364211,0.00074745616,0.00026563244,0.0005191708],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00548688,0.00057486026,0.14187992,0.00050527096,0.0007225995,0.004767177,0.0014498436,0.044111755,0.7069005,0.029615331,0.004283415,0.05970246],"study_design_scores_gemma":[0.00069051405,0.0022167442,0.19359753,0.000039361683,0.00069069356,0.0031621756,0.0010346418,0.36463434,0.40412965,0.020186901,0.009295082,0.00032244597],"about_ca_topic_score_codex":0.0027731787,"about_ca_topic_score_gemma":0.00376411,"teacher_disagreement_score":0.0027731787,"about_ca_system_score_codex":0.0005253995,"about_ca_system_score_gemma":0.0005032015,"threshold_uncertainty_score":0.007535517},"labels":[],"label_agreement":null},{"id":"W7024823872","doi":"","title":"The 2016 Stubbendieck Great Plains Distinguished Book Prize Winner: Michel Hogue’s <i>Metis and the Medicine Line: Creating a Border and Dividing a People</i>","year":2016,"lang":"en","type":"article","venue":"Project Muse (Johns Hopkins University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Government (linguistics); Feature (linguistics); Agency (philosophy); Work (physics)","score_opus":0.012682927566012256,"score_gpt":0.2249796667395077,"score_spread":0.21229673917349545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024823872","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018833958,0.02192609,0.0043853642,0.5605652,0.32155088,0.00009487403,0.0010926669,0.00050954655,0.08799199],"genre_scores_gemma":[0.014748043,0.0068193837,0.0030626184,0.029918224,0.04090304,0.00010191422,0.0010405534,0.0009661397,0.9024401],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981589,0.00044699357,0.000053190975,0.00029264245,0.00079381856,0.00025451396],"domain_scores_gemma":[0.9957847,0.0005253249,0.00012166192,0.00017557031,0.0007795363,0.0026131724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004193314,0.0010986552,0.0009675795,0.0017759604,0.0036875168,0.014416823,0.0011666128,0.0028216213,0.034360517],"category_scores_gemma":[0.008378101,0.0004640009,0.00082333316,0.001496159,0.002223208,0.004712291,0.0047841663,0.006030915,0.019205736],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000088429215,0.0000041737076,0.000052374417,0.000013156578,0.000002933329,0.000009593995,0.000044443186,0.000024429062,0.000030407737,0.0031567973,0.99136823,0.00528466],"study_design_scores_gemma":[0.000005441495,0.0000073577166,0.00029775986,0.0000513856,0.0000034401787,0.000027868253,0.0003201035,0.00010102664,0.00006412526,0.0036238187,0.9954893,0.000008404628],"about_ca_topic_score_codex":0.01076137,"about_ca_topic_score_gemma":0.054697484,"teacher_disagreement_score":0.9892386,"about_ca_system_score_codex":0.0046315994,"about_ca_system_score_gemma":0.005058299,"threshold_uncertainty_score":0.11494744},"labels":[],"label_agreement":null},{"id":"W7025041631","doi":"","title":"Tenir l'Ã©vanouissement : entre maÃ®trise intÃ©grale et abandon anÃ©antissant : Jean Genet et Antonin Artaud","year":2012,"lang":"fr","type":"other","venue":"Library and Archives Canada (Government of Canada)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Chose; Ostensive definition; Intentionality","score_opus":0.006010944715226579,"score_gpt":0.16004459727986575,"score_spread":0.15403365256463916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7025041631","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031259358,0.07703753,0.00640455,0.59551805,0.028192028,0.000045585977,0.00006081467,0.00029522477,0.26118675],"genre_scores_gemma":[0.23519385,0.021424852,0.0032440757,0.057750117,0.0052727074,0.00006482996,0.00006116841,0.0007441878,0.67624414],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.994288,0.0024188815,0.0001141361,0.0008835282,0.001477326,0.00081816554],"domain_scores_gemma":[0.99546695,0.0014809732,0.00033032818,0.00024911496,0.0012355334,0.0012370426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050486946,0.00065483514,0.00052761444,0.00083130004,0.012258526,0.015304349,0.0013791476,0.0048425407,0.0150299305],"category_scores_gemma":[0.008696089,0.00043406867,0.00045702548,0.00079941296,0.011995054,0.008580057,0.0054289536,0.009955665,0.0045226617],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015270652,0.00007521364,0.0021655024,0.00016624873,0.000020990548,0.001194792,0.060342044,0.0002138413,0.0011018886,0.42555758,0.42963403,0.079375125],"study_design_scores_gemma":[0.000003163228,0.000013554684,0.00043308502,0.00011441877,0.0000029916098,0.00024077564,0.006975784,0.00009944032,0.0002370206,0.008867961,0.98299545,0.000016354623],"about_ca_topic_score_codex":0.023867523,"about_ca_topic_score_gemma":0.03614405,"teacher_disagreement_score":0.023867523,"about_ca_system_score_codex":0.007664979,"about_ca_system_score_gemma":0.008933474,"threshold_uncertainty_score":0.055613577},"labels":[],"label_agreement":null},{"id":"W7025063110","doi":"","title":"Travailleurs (im)migrants au Québec et au Canada : vers le respect administratif de leurs droits et libertés ?","year":2008,"lang":"fr","type":"report","venue":"Érudit documents and data repository (Érudit Consortium, University of Montreal)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Nucleofection; TSG101; Pretext; Gestational period; Liquation; Hyporeflexia","score_opus":0.04765167940614926,"score_gpt":0.25519846816046005,"score_spread":0.2075467887543108,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7025063110","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85314536,0.002965873,0.0010310125,0.0139240455,0.00009154343,0.00006472733,0.00046186964,0.00002902802,0.12828651],"genre_scores_gemma":[0.9686073,0.0008578883,0.0003484795,0.0004944916,0.000009436889,0.000017151771,0.00008114121,0.000009460894,0.029574549],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988702,0.00017529544,0.000019111532,0.0001128897,0.00022844005,0.00059407135],"domain_scores_gemma":[0.99870443,0.0001033138,0.00016076442,0.000045745383,0.00050162553,0.00048419228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00092322065,0.00024524037,0.00025674512,0.0007591524,0.007818163,0.005048709,0.0007923431,0.00069662824,0.0060537593],"category_scores_gemma":[0.0016108787,0.00012368413,0.0002217893,0.0017256374,0.0050060977,0.0012285063,0.0015522446,0.0012614015,0.00024001984],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003225223,0.0001884071,0.28337744,0.00030503314,0.00008401033,0.0013697069,0.2188605,0.0027361047,0.0025175235,0.3121933,0.023690928,0.15435451],"study_design_scores_gemma":[0.00003491374,0.00014065436,0.47310844,0.0006765312,0.00006482067,0.0001832688,0.19534793,0.0011132653,0.00078871514,0.0061068954,0.32232282,0.00011177033],"about_ca_topic_score_codex":0.99293774,"about_ca_topic_score_gemma":0.9968771,"teacher_disagreement_score":0.08579115,"about_ca_system_score_codex":0.08579115,"about_ca_system_score_gemma":0.10522454,"threshold_uncertainty_score":0.62246126},"labels":[],"label_agreement":null},{"id":"W7025080835","doi":"","title":"Translation into English: Poems by Marie-Léontine Tsibinda","year":2012,"lang":"en","type":"article","venue":"UNI ScholarWorks (University of Northern Iowa)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Poetry; Context (archaeology); Homeland; Spanish Civil War; Capital (architecture); French","score_opus":0.01576878406979601,"score_gpt":0.19443011695503196,"score_spread":0.17866133288523595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7025080835","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07617212,0.094830565,0.0071950178,0.2041846,0.051382374,0.00009392031,0.0006196861,0.00032135006,0.5652004],"genre_scores_gemma":[0.71151686,0.03406394,0.003114731,0.013132911,0.0077149537,0.0001126761,0.00040745034,0.00067422894,0.2292622],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99950147,0.0003155133,0.000014910993,0.00005076755,0.00006673994,0.000050636703],"domain_scores_gemma":[0.99943167,0.00037075244,0.000035306814,0.000023730563,0.00008810822,0.000050366267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008755195,0.00063981494,0.0002599727,0.00047094448,0.0033717381,0.0022307732,0.00027979744,0.0006655656,0.005082462],"category_scores_gemma":[0.0034075254,0.00014010795,0.000121303485,0.00051422475,0.002741317,0.0019001575,0.0012483348,0.002203158,0.0015521949],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001441515,0.000049976596,0.00083203113,0.00045221264,0.000017057326,0.003040325,0.16832137,0.00041320463,0.0022665826,0.16305728,0.6061495,0.055256236],"study_design_scores_gemma":[0.0000031745294,0.000009335285,0.000499456,0.00018421611,0.000002357767,0.0005859568,0.010169447,0.00012296779,0.0002055199,0.0018382977,0.98637193,0.000007343705],"about_ca_topic_score_codex":0.008160408,"about_ca_topic_score_gemma":0.01358964,"teacher_disagreement_score":0.008160408,"about_ca_system_score_codex":0.0019040711,"about_ca_system_score_gemma":0.00075427303,"threshold_uncertainty_score":0.017002523},"labels":[],"label_agreement":null},{"id":"W7025275728","doi":"","title":"Use of foliar calcium to strontium ratios to partition soil calcium sources of American beech on two sites in Southern Québec","year":2011,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Beech; Strontium; Calcium; Partition (number theory); Soil water; Sink (geography)","score_opus":0.04998682501051842,"score_gpt":0.27605094631688554,"score_spread":0.2260641213063671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7025275728","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9980286,0.0000842464,0.0002165436,0.00003583567,0.0000021951796,0.000033570363,0.0005148669,0.000011269195,0.0010728879],"genre_scores_gemma":[0.9970471,0.00008781438,0.0006342232,0.00002435547,0.0000014644477,0.000023007939,0.00045916214,0.000009690728,0.0017132113],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99979097,0.0000327918,0.000006617418,0.000072104995,0.00004260866,0.000054854692],"domain_scores_gemma":[0.99929047,0.00015306466,0.00004966946,0.000013869159,0.00038252352,0.00011040626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0003907707,0.00037622234,0.0004341392,0.0011001428,0.002664153,0.0012691185,0.0009018995,0.0005471668,0.0013419247],"category_scores_gemma":[0.00068235723,0.00025580978,0.00027715997,0.0012873686,0.0005814108,0.00030268295,0.00036586614,0.00036973323,0.00023131295],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00081502786,0.0002223909,0.9284197,0.000117936914,0.00021716424,0.0005055985,0.0067449883,0.0022971123,0.035871994,0.0001845682,0.0010289004,0.02357472],"study_design_scores_gemma":[0.00001519397,0.00003684578,0.9943402,0.000012196155,0.000032345833,0.000020163401,0.0018871858,0.0020846755,0.0008565785,0.000023041805,0.00068001857,0.0000115867315],"about_ca_topic_score_codex":0.97778416,"about_ca_topic_score_gemma":0.99372995,"teacher_disagreement_score":0.022215843,"about_ca_system_score_codex":0.009619936,"about_ca_system_score_gemma":0.0036842227,"threshold_uncertainty_score":0.06979787},"labels":[],"label_agreement":null},{"id":"W7026899228","doi":"","title":"Autoencoders for natural language semantics","year":2022,"lang":"fr","type":"dissertation","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Compute Canada","keywords":"Formal semantics (linguistics); Domain (mathematical analysis); Semantic interpretation","score_opus":0.0082578758154122,"score_gpt":0.19891413347091397,"score_spread":0.19065625765550176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7026899228","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046050884,0.0032349296,0.97620684,0.0009916103,0.0003185492,0.00009041317,0.0017803307,0.007471005,0.0053012935],"genre_scores_gemma":[0.21507864,0.008144989,0.7307681,0.0011014349,0.0007235424,0.00069104764,0.009730081,0.0021171442,0.03164506],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99890697,0.00029049683,0.00011199479,0.0003202137,0.00029810244,0.000072264185],"domain_scores_gemma":[0.998209,0.0011223796,0.00011427746,0.00026920118,0.0002504409,0.00003475819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012583607,0.0011991789,0.0009912807,0.0016757248,0.00056032185,0.0023099105,0.0015114421,0.001217735,0.014243427],"category_scores_gemma":[0.0045847455,0.0009269994,0.0017063316,0.0014872674,0.0010360216,0.003974861,0.0014717879,0.002916169,0.005206999],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002767523,0.000134994,0.0010156811,0.0008656638,0.00026054168,0.0002802593,0.00032373337,0.112174794,0.006414336,0.19884561,0.030426705,0.64898086],"study_design_scores_gemma":[0.000041622246,0.0000487428,0.0006203871,0.00017115123,0.0000368488,0.00010917011,0.00006981359,0.6940964,0.0029307571,0.26616275,0.035679776,0.00003265728],"about_ca_topic_score_codex":0.008219033,"about_ca_topic_score_gemma":0.0149421925,"teacher_disagreement_score":0.014243427,"about_ca_system_score_codex":0.0013856726,"about_ca_system_score_gemma":0.0015533454,"threshold_uncertainty_score":0.047649026},"labels":[],"label_agreement":null},{"id":"W7029929027","doi":"","title":"Leveraging lexical resources as external knowledge for entity reasoning using deep learning frameworks","year":2017,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Leverage (statistics); Knowledge base; Question answering; Taxonomy (biology); Natural language; Deep learning","score_opus":0.030777405744136344,"score_gpt":0.29940122092200516,"score_spread":0.2686238151778688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7029929027","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.058131788,0.0008949396,0.93287057,0.0013720972,0.0000886258,0.000096602416,0.00037234742,0.002100259,0.0040728273],"genre_scores_gemma":[0.6328622,0.0009696096,0.35572684,0.0005887806,0.00012020073,0.00021535714,0.001708797,0.0001752069,0.007633134],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994242,0.0001610206,0.000050613675,0.00018095608,0.00011213296,0.000071131006],"domain_scores_gemma":[0.9982962,0.0009284214,0.0001443412,0.0002604813,0.00027987332,0.000090653666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015386693,0.00089654454,0.00069086434,0.0012775139,0.00066470855,0.0023768942,0.0020794235,0.001431672,0.0032603242],"category_scores_gemma":[0.0047484953,0.00066680973,0.0010993576,0.0009869325,0.00085370924,0.0061799674,0.0023632431,0.002517604,0.0011502195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022670905,0.0004134717,0.005684743,0.00033844364,0.0002535737,0.0005402777,0.0006911766,0.4719108,0.012582141,0.0448544,0.0072511304,0.4552531],"study_design_scores_gemma":[0.0000053870267,0.000015117543,0.00020264101,0.000024891502,0.000021227086,0.00002252271,0.000043878514,0.9783524,0.0016645959,0.01819003,0.0014495399,0.00000777055],"about_ca_topic_score_codex":0.008272056,"about_ca_topic_score_gemma":0.0159336,"teacher_disagreement_score":0.008272056,"about_ca_system_score_codex":0.0011797223,"about_ca_system_score_gemma":0.0013544855,"threshold_uncertainty_score":0.016447783},"labels":[],"label_agreement":null},{"id":"W7033532016","doi":"","title":"\" Protecting Mauna Kea from the $1.4 Billion Thirty Meter Telescope (TMT) Project\" - March 2nd, 2019","year":2019,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Indigenous; Observatory; Telescope; Government (linguistics); Research council; Geological survey","score_opus":0.01135025571396783,"score_gpt":0.20065597356481707,"score_spread":0.18930571785084924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7033532016","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018679278,0.013119859,0.0015715831,0.20959798,0.0224825,0.00064749795,0.0047370475,0.00083422696,0.72833014],"genre_scores_gemma":[0.028244648,0.0027828256,0.00090315007,0.030114567,0.00094384264,0.000269846,0.0015582266,0.0001646791,0.9350182],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99941874,0.000067386354,0.000021356347,0.000059559912,0.00020818887,0.00022493117],"domain_scores_gemma":[0.9992447,0.0000544098,0.00004249941,0.000040046867,0.0003029692,0.00031541087],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009466462,0.0005050852,0.0001441203,0.0003581545,0.0050884574,0.0035287864,0.00076725875,0.0047464697,0.0509114],"category_scores_gemma":[0.0017523841,0.0002555855,0.000271574,0.00029538965,0.0010307403,0.0016631063,0.0027958266,0.004837226,0.018253176],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025763564,0.00002690601,0.0015421909,0.000112712136,0.0000049615196,0.00024903755,0.0013961276,0.000030469722,0.0010340225,0.0039813793,0.9762279,0.015368561],"study_design_scores_gemma":[0.0000047565336,0.000017167267,0.0029097497,0.00010097128,0.0000034946074,0.000038890124,0.0020870182,0.000017001343,0.0002559349,0.0002313216,0.9943243,0.000009339345],"about_ca_topic_score_codex":0.13159247,"about_ca_topic_score_gemma":0.30462193,"teacher_disagreement_score":0.13159247,"about_ca_system_score_codex":0.0021106205,"about_ca_system_score_gemma":0.00832795,"threshold_uncertainty_score":0.261653},"labels":[],"label_agreement":null},{"id":"W7033610217","doi":"","title":"Representació de la llegenda de Sant Jordi","year":2011,"lang":"ca","type":"other","venue":"Repositori UJI (Universitat Jaume I)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Quarter (Canadian coin); Population; Context (archaeology)","score_opus":0.013048300078376485,"score_gpt":0.23813541085812198,"score_spread":0.2250871107797455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7033610217","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18206137,0.00989048,0.037281048,0.027122201,0.002733589,0.00048918976,0.009558822,0.003572801,0.7272906],"genre_scores_gemma":[0.7318417,0.0047188336,0.051126737,0.001345666,0.00044215843,0.00047289333,0.0099501405,0.0009460031,0.1991559],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977191,0.0008197793,0.00007289297,0.0004525238,0.00063585385,0.0002998601],"domain_scores_gemma":[0.9953773,0.0016983176,0.0004139016,0.00048141394,0.0011233333,0.00090573216],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034067738,0.00035217963,0.00030491606,0.0025489938,0.0026558419,0.006068657,0.0008276729,0.0012058492,0.020141875],"category_scores_gemma":[0.007201447,0.00033177107,0.00056768407,0.0031774985,0.0012015279,0.0024202357,0.0028605615,0.0014320558,0.0037032412],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063763076,0.00017368913,0.05185793,0.0011567508,0.000071687944,0.000830804,0.06000757,0.004356337,0.005277848,0.12388793,0.40632725,0.3454146],"study_design_scores_gemma":[0.000013830713,0.000052354004,0.022336591,0.00032596829,0.000026348756,0.00024230097,0.008839623,0.0024047303,0.0010147744,0.0048330342,0.9598662,0.000044157518],"about_ca_topic_score_codex":0.057418548,"about_ca_topic_score_gemma":0.085965365,"teacher_disagreement_score":0.057418548,"about_ca_system_score_codex":0.005545232,"about_ca_system_score_gemma":0.0059967106,"threshold_uncertainty_score":0.1141687},"labels":[],"label_agreement":null},{"id":"W7035772564","doi":"","title":"ABA) to Present at Canadian Marijuana Investor Conference Held by Jacob Securities","year":2015,"lang":"en","type":"other","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Government (linguistics); Legislation; Feature (linguistics); Payment","score_opus":0.03482790113155616,"score_gpt":0.2356720479290123,"score_spread":0.20084414679745616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7035772564","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015957042,0.0015383497,0.0048913024,0.0101146735,0.005889499,0.00015299059,0.008011483,0.0035079373,0.964298],"genre_scores_gemma":[0.003117987,0.00036418872,0.00074020866,0.00045492125,0.00036844335,0.000021588448,0.001137922,0.00028360152,0.99351114],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997212,0.000023706492,0.0000074301847,0.000046418907,0.0001224934,0.00007871244],"domain_scores_gemma":[0.99868256,0.00010189975,0.000049680824,0.000082625345,0.0005171897,0.0005660862],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00077891536,0.00073405745,0.0005718229,0.0014268061,0.0019414963,0.0035639498,0.0009802254,0.0012752194,0.67708814],"category_scores_gemma":[0.0019902256,0.00024759967,0.00042908412,0.0012846757,0.00032914296,0.0020720756,0.0014755123,0.0012336829,0.4306939],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031505988,0.000043048676,0.00026992522,0.00003329365,0.000002865671,0.000023625018,0.00003418899,0.000082363695,0.00017340742,0.00217962,0.92667526,0.070450984],"study_design_scores_gemma":[0.000009012481,0.000012200588,0.00081633416,0.00003469279,0.0000047368685,0.000013045604,0.000077595934,0.0003826442,0.0002024488,0.0019183066,0.9965228,0.0000062296085],"about_ca_topic_score_codex":0.053627845,"about_ca_topic_score_gemma":0.25515723,"teacher_disagreement_score":0.94637215,"about_ca_system_score_codex":0.0019989433,"about_ca_system_score_gemma":0.0033596077,"threshold_uncertainty_score":0.46059453},"labels":[],"label_agreement":null},{"id":"W7039589344","doi":"","title":"A millennial housing typology: redefining the home of the next generation in Toronto laneways","year":2019,"lang":"en","type":"dissertation","venue":"Lu Zone Ul (Laurentian University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Generation y; Architecture; Generation x; Fourth generation; Term (time)","score_opus":0.02491211291731056,"score_gpt":0.21937915962141183,"score_spread":0.19446704670410125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039589344","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46683145,0.017530307,0.0041384576,0.022454426,0.00125072,0.00022368557,0.0021899154,0.00008419345,0.48529693],"genre_scores_gemma":[0.9541576,0.0055206968,0.0018219082,0.00082968903,0.00006500164,0.000081201244,0.00048776512,0.000056912995,0.036979187],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9993218,0.00017452463,0.000027189451,0.000058978814,0.00012909854,0.00028833994],"domain_scores_gemma":[0.99925846,0.000056285844,0.000060947677,0.000040060448,0.00018053173,0.00040364004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005705264,0.0003483991,0.00016047143,0.0016158298,0.012622464,0.005059905,0.0012714578,0.0008257401,0.008572937],"category_scores_gemma":[0.0006562609,0.00021848422,0.00026636536,0.002881229,0.0075281956,0.002964696,0.0048778383,0.0013486332,0.0005104625],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044031563,0.000033849406,0.032741617,0.0002352738,0.000013357442,0.0024243935,0.5413528,0.00056471466,0.00043879333,0.33799285,0.05520963,0.028948633],"study_design_scores_gemma":[0.0000038668236,0.000020251471,0.04361581,0.0005400632,0.000013291466,0.000520827,0.55627203,0.00028442862,0.00014552416,0.005397,0.39315665,0.000030246403],"about_ca_topic_score_codex":0.9194328,"about_ca_topic_score_gemma":0.97341996,"teacher_disagreement_score":0.08056718,"about_ca_system_score_codex":0.04661106,"about_ca_system_score_gemma":0.021912877,"threshold_uncertainty_score":0.33818847},"labels":[],"label_agreement":null},{"id":"W7039643056","doi":"","title":"Microbial diversity of buckwheat rhizosphere in wireworm-infested and non-infested soils using metagenomics","year":2019,"lang":"en","type":"article","venue":"IslandScholar (University of Prince Edward Island)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Rhizosphere; Metagenomics; Crop rotation; Context (archaeology); Crop; Actinobacteria; Microbial population biology; Population; Soil water","score_opus":0.010846166192874996,"score_gpt":0.19988007173931704,"score_spread":0.18903390554644206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039643056","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98432565,0.0019496245,0.001998423,0.00012592235,0.000017173563,0.00006422493,0.010528293,0.00006259379,0.0009281378],"genre_scores_gemma":[0.97215265,0.0022021777,0.0073278383,0.00028047603,0.000023845121,0.00018082664,0.015659677,0.000052994114,0.0021195048],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9995641,0.000039325107,0.000027371678,0.0001930079,0.00009109017,0.000085119325],"domain_scores_gemma":[0.9997687,0.00003895849,0.00005437923,0.000012651564,0.00009423699,0.000031115975],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00035727807,0.00075395685,0.0008199384,0.0010991289,0.00053878356,0.0012135244,0.00027041935,0.00060844637,0.00076842174],"category_scores_gemma":[0.00040591322,0.00025026454,0.00073155115,0.0013825322,0.0002541174,0.0006335308,0.00062936853,0.0005686656,0.00036285096],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085057865,0.00019875876,0.08709981,0.0008921131,0.00029511217,0.0003322302,0.0012127653,0.0007339397,0.8869271,0.00013709444,0.00039431851,0.020926118],"study_design_scores_gemma":[0.000027603546,0.0008347678,0.90391433,0.00019793032,0.0004458836,0.00062891387,0.0028186396,0.003482202,0.07901837,0.0004365888,0.00810394,0.00009084122],"about_ca_topic_score_codex":0.004683518,"about_ca_topic_score_gemma":0.007677001,"teacher_disagreement_score":0.004683518,"about_ca_system_score_codex":0.00035335962,"about_ca_system_score_gemma":0.0005140459,"threshold_uncertainty_score":0.0093125105},"labels":[],"label_agreement":null},{"id":"W7039777345","doi":"","title":"Migration policy in a small open economy with a dual labor market","year":2003,"lang":"en","type":"article","venue":"Archive ouverte UNIGE (University of Geneva)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Université Laval; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; National Science Foundation","keywords":"Small open economy; Dual (grammatical number); Secondary labor market; Margin (machine learning); Factor market","score_opus":0.01422662267055344,"score_gpt":0.1936045554669704,"score_spread":0.17937793279641695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039777345","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97310424,0.00034231838,0.009041654,0.005833882,0.0000839385,0.00004110546,0.0002794105,0.00007542239,0.011198063],"genre_scores_gemma":[0.99328494,0.00015100367,0.000685423,0.000078014986,0.00004266869,0.000018171499,0.00004471585,0.0000074210025,0.005687808],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9994722,0.00024694056,0.000022399667,0.00007157551,0.000026632017,0.00016016973],"domain_scores_gemma":[0.9967894,0.0017363676,0.0004026836,0.0001291423,0.0001759794,0.00076636276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012139371,0.00032643919,0.0011049244,0.00063691556,0.0016857313,0.002881892,0.00094981503,0.0018866844,0.007893052],"category_scores_gemma":[0.003998861,0.0003485082,0.0004465558,0.0007231847,0.0014907563,0.0021310933,0.001493522,0.0012618377,0.0005233515],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039538657,0.0016573514,0.022958482,0.00046326083,0.00017115888,0.003936065,0.0018815587,0.40743724,0.004016383,0.49009752,0.021690449,0.0417367],"study_design_scores_gemma":[0.0009604579,0.0005600306,0.012401244,0.00007187189,0.00016972024,0.00023464632,0.005586674,0.6288226,0.0007832546,0.3433716,0.006961364,0.0000765108],"about_ca_topic_score_codex":0.01261005,"about_ca_topic_score_gemma":0.01237269,"teacher_disagreement_score":0.01261005,"about_ca_system_score_codex":0.001346295,"about_ca_system_score_gemma":0.001419399,"threshold_uncertainty_score":0.026404917},"labels":[],"label_agreement":null},{"id":"W7066382463","doi":"","title":"Information Retrieval with Dense and Sparse Representations","year":2024,"lang":"en","type":"dissertation","venue":"DSpace@MIT (Massachusetts Institute of Technology)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Query expansion; Relevance (law); Question answering; Pipeline (software); Language model; Matching (statistics); Bottleneck; Representation (politics); Sentence","score_opus":0.011760574607934884,"score_gpt":0.2525750935037714,"score_spread":0.2408145188958365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7066382463","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012571712,0.0014852239,0.98103255,0.0010230141,0.00007632292,0.00012251321,0.0003885961,0.0008374678,0.0024626064],"genre_scores_gemma":[0.28873634,0.0035279556,0.69439805,0.0008403029,0.0005657482,0.00043304573,0.0028590288,0.0002448333,0.0083948225],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982015,0.0007252051,0.00013511794,0.00039043554,0.00042093595,0.00012695845],"domain_scores_gemma":[0.9960205,0.0021856087,0.0001924574,0.0010677024,0.00044478668,0.00008892731],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002208291,0.0006900155,0.001190749,0.0019352186,0.0005017842,0.0023723543,0.0013112094,0.0011810461,0.0030730097],"category_scores_gemma":[0.010634735,0.0005021651,0.0012763689,0.002385873,0.0011497408,0.006961336,0.0025549326,0.0021225682,0.0018898417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028037804,0.00027786018,0.00117756,0.00088624377,0.00016360843,0.00023540521,0.00073099526,0.11303609,0.022497775,0.20882435,0.024529906,0.62735987],"study_design_scores_gemma":[0.000051698407,0.00011848089,0.00045372394,0.00005610472,0.000050162416,0.00017374374,0.00014509913,0.76201564,0.0057662316,0.21805893,0.01306801,0.000042200165],"about_ca_topic_score_codex":0.0020031102,"about_ca_topic_score_gemma":0.0023894305,"teacher_disagreement_score":0.0030730097,"about_ca_system_score_codex":0.0010441542,"about_ca_system_score_gemma":0.00094685005,"threshold_uncertainty_score":0.011678696},"labels":[],"label_agreement":null},{"id":"W7067080718","doi":"","title":"Know Your Oil: Creating A Global Oil-Climate Index","year":2015,"lang":"en","type":"report","venue":"Issue Lab (Candid)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Midstream; Index (typography); Upstream (networking); Climate change; Greenhouse gas; Tonne; Global warming; Downstream (manufacturing); Fossil fuel","score_opus":0.046338010965858684,"score_gpt":0.32784205975726766,"score_spread":0.281504048791409,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7067080718","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13302447,0.0034741121,0.26061186,0.016986046,0.0023717121,0.0014177925,0.13754845,0.017550727,0.42701474],"genre_scores_gemma":[0.33998036,0.0044863266,0.42452487,0.0012056988,0.0010056032,0.0015997803,0.18233547,0.0042979163,0.04056409],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99864286,0.00031608305,0.00011110149,0.00018277318,0.0006205259,0.0001267711],"domain_scores_gemma":[0.99491286,0.0016000235,0.00040890853,0.00065950287,0.0017985543,0.00062004157],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003616022,0.0006968563,0.00035655452,0.007839427,0.0006347271,0.0029814593,0.0005523115,0.00056224485,0.011521778],"category_scores_gemma":[0.012665009,0.00026559696,0.00041880624,0.010093411,0.000291524,0.006222584,0.002493425,0.0008960919,0.0045098932],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012999661,0.00020258637,0.041199673,0.00031594915,0.000075715216,0.00010094914,0.00061942136,0.012510427,0.0023136565,0.054284148,0.29963464,0.58861285],"study_design_scores_gemma":[0.000032928307,0.00010660788,0.027324,0.00013824255,0.00005712685,0.0001294782,0.0013561229,0.031217134,0.0045786304,0.040847056,0.89413565,0.00007705534],"about_ca_topic_score_codex":0.0054962393,"about_ca_topic_score_gemma":0.008004187,"teacher_disagreement_score":0.011521778,"about_ca_system_score_codex":0.0010421874,"about_ca_system_score_gemma":0.0012545341,"threshold_uncertainty_score":0.038544238},"labels":[],"label_agreement":null},{"id":"W7077911890","doi":"10.48448/7ydn-np56","title":"ResearchAgent: Iterative Research Idea Generation over Scientific Literature with Large Language Models","year":2025,"lang":"en","type":"other","venue":"Underline Science Inc.","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Operationalization; Leverage (statistics); Pace; Language model; Scientific literature; Readability; Natural language generation","score_opus":0.06560478267328029,"score_gpt":0.37216658307195777,"score_spread":0.3065618003986775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7077911890","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02161478,0.0005467556,0.9021878,0.0012024523,0.00021692693,0.0012432487,0.0008482898,0.06870091,0.0034389533],"genre_scores_gemma":[0.11899083,0.00028823206,0.8707131,0.00051274686,0.00013989687,0.0012566374,0.0024747597,0.0019296723,0.0036941231],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9876881,0.0077993334,0.0008343005,0.001700756,0.0017702681,0.00020719171],"domain_scores_gemma":[0.93140066,0.049668487,0.0033497019,0.009381078,0.0045742355,0.0016257125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017552026,0.002200178,0.0012563458,0.0035156286,0.0011398048,0.0046333037,0.0040550157,0.0023192964,0.007430641],"category_scores_gemma":[0.06767564,0.0013219663,0.0022699449,0.0016539847,0.0014019769,0.008203057,0.0071644136,0.0021946053,0.004665955],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019608499,0.0015087536,0.01123626,0.002947437,0.000926972,0.0010948598,0.007937774,0.05039306,0.037306152,0.043405067,0.066760704,0.774522],"study_design_scores_gemma":[0.0006960353,0.0006232221,0.0013557315,0.00025330595,0.0003422895,0.00041148762,0.0011296668,0.8302452,0.024963843,0.058184553,0.08156812,0.00022661527],"about_ca_topic_score_codex":0.0018834396,"about_ca_topic_score_gemma":0.004196798,"teacher_disagreement_score":0.017552026,"about_ca_system_score_codex":0.0013326756,"about_ca_system_score_gemma":0.0030171315,"threshold_uncertainty_score":0.092825115},"labels":[],"label_agreement":null},{"id":"W7082252666","doi":"10.48448/jjkr-0706","title":"Navigating the Prompt Space: Supervision Matters in CoT When Reasoning Misleads","year":2025,"lang":"en","type":"other","venue":"Underline Science Inc.","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Task (project management); Forcing (mathematics); Function (biology); Architecture; Foundation (evidence); State (computer science); Trajectory; SAFER","score_opus":0.015896805927801207,"score_gpt":0.2887288984631789,"score_spread":0.2728320925353777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7082252666","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20418456,0.00076231075,0.7707949,0.003569066,0.00016759548,0.0002669519,0.0004614583,0.01093218,0.008861072],"genre_scores_gemma":[0.787102,0.00025703313,0.20832509,0.00087984797,0.00005078331,0.00013656163,0.0006255375,0.00087597826,0.0017472137],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9938567,0.0025160355,0.00044390553,0.001649169,0.0011221436,0.00041203597],"domain_scores_gemma":[0.94150317,0.0403365,0.003752666,0.010060481,0.0027609456,0.0015862097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008752737,0.0012645537,0.0011633055,0.0006792438,0.0011816365,0.0042492067,0.0020136426,0.002529169,0.0058996705],"category_scores_gemma":[0.090962335,0.0009858523,0.0011873336,0.0006696127,0.0029160038,0.011607981,0.004480085,0.005399108,0.0014899053],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038424425,0.0013076544,0.0658092,0.0018754799,0.00024388298,0.0018018641,0.016310045,0.113531284,0.051890317,0.117531665,0.01783015,0.608026],"study_design_scores_gemma":[0.0002993108,0.0006295055,0.0059538377,0.00027455328,0.00017695583,0.0008754141,0.0025842444,0.65699726,0.024854213,0.29414538,0.013063838,0.00014548788],"about_ca_topic_score_codex":0.004246028,"about_ca_topic_score_gemma":0.0044594714,"teacher_disagreement_score":0.008752737,"about_ca_system_score_codex":0.0014763142,"about_ca_system_score_gemma":0.0035476123,"threshold_uncertainty_score":0.046289444},"labels":[],"label_agreement":null},{"id":"W7084070425","doi":"10.6084/m9.figshare.c.8058552.v1","title":"Coenrollment of critically ill patients in PROSPECT: characteristics and association with treatment efficacy and safety","year":2025,"lang":"en","type":"other","venue":"Figshare","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Joseph Brant Hospital; University of Manitoba; Western University; Hôpital Maisonneuve-Rosemont; McMaster University; Queen's University; University of Ottawa; Université Laval; University of Toronto","funders":"","keywords":"Adverse effect; Randomized controlled trial; Logistic regression; Pneumonia; Critically ill; Clinical trial; Proportional hazards model; Clinical endpoint","score_opus":0.011458159006560801,"score_gpt":0.23133512736218884,"score_spread":0.21987696835562803,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7084070425","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9725705,0.019570487,0.0029622524,0.00069823605,0.000109644745,0.00087399705,0.001824019,0.000024436555,0.0013663513],"genre_scores_gemma":[0.996566,0.00081839977,0.0012289122,0.00021995245,0.00006243323,0.00049787795,0.0004760429,0.0000061437668,0.00012426094],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9675225,0.01588267,0.008565049,0.0036139018,0.003628247,0.0007875283],"domain_scores_gemma":[0.900947,0.0387138,0.050250813,0.005433511,0.0029488879,0.0017060377],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023720315,0.0004838283,0.0014237623,0.0019569779,0.0006597457,0.0017265733,0.00082249567,0.0010192338,0.0038112048],"category_scores_gemma":[0.06372911,0.00030939284,0.0035572674,0.0022931013,0.0010368392,0.001255259,0.001712493,0.0008794122,0.00019196594],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.03518191,0.00035434047,0.9002049,0.00492931,0.018046534,0.00024093599,0.0006896105,0.00047675968,0.0025572025,0.00067468046,0.0010040958,0.035639748],"study_design_scores_gemma":[0.003275115,0.008258477,0.9646171,0.0017248658,0.011013766,0.001288647,0.00073677656,0.0013738162,0.0021917738,0.0014609053,0.003951532,0.00010728514],"about_ca_topic_score_codex":0.0004321589,"about_ca_topic_score_gemma":0.00083051895,"teacher_disagreement_score":0.023720315,"about_ca_system_score_codex":0.00056644523,"about_ca_system_score_gemma":0.00085791666,"threshold_uncertainty_score":0.1254465},"labels":[],"label_agreement":null},{"id":"W7084071777","doi":"10.64628/aap.k36yjax6t","title":"The White Lotus en Thaïlande : l’envers du tourisme de luxe","year":2025,"lang":"fr","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"White (mutation); Patron saint; Key (lock)","score_opus":0.011041744362331143,"score_gpt":0.23432771019522597,"score_spread":0.22328596583289484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7084071777","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.904216,0.0055113356,0.0060145417,0.012530189,0.00035954206,0.000070825714,0.002844208,0.00018802786,0.06826524],"genre_scores_gemma":[0.9244248,0.0022937853,0.002627963,0.00042884075,0.00014552008,0.000053768188,0.0014443268,0.000118493306,0.068462424],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996803,0.00013021634,0.000008325929,0.000049399103,0.000062947496,0.00006885112],"domain_scores_gemma":[0.9993538,0.000252115,0.00006451868,0.000021316035,0.00009660283,0.00021152927],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005947897,0.00036759043,0.0001814648,0.00088167505,0.0024860555,0.0038147478,0.00039933794,0.0007794459,0.006701943],"category_scores_gemma":[0.0012820538,0.00020648527,0.00024847494,0.0024258764,0.0010042704,0.0020964723,0.0012210488,0.0011662574,0.0006726348],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017513015,0.0004590513,0.27888024,0.0011689648,0.000293576,0.01559729,0.19488582,0.022948768,0.0048377425,0.09081443,0.14495754,0.24340522],"study_design_scores_gemma":[0.000075598844,0.00024181008,0.2094498,0.0005849112,0.00012367571,0.0023077377,0.20209935,0.021834059,0.001981047,0.009583596,0.5515666,0.0001518033],"about_ca_topic_score_codex":0.20434724,"about_ca_topic_score_gemma":0.3000668,"teacher_disagreement_score":0.20434724,"about_ca_system_score_codex":0.002469411,"about_ca_system_score_gemma":0.0024369252,"threshold_uncertainty_score":0.40631562},"labels":[],"label_agreement":null},{"id":"W7084253775","doi":"","title":"Comparison of sperm availability in hybridizing male Fundulus heteroclitus and F. diaphanus in Porters Lake, Nova Scotia","year":2025,"lang":"en","type":"article","venue":"Saint Mary's University Institutional Repository (Saint Mary's University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Fundulus; Nova scotia; Sperm; Resource (disambiguation); Population; Gastropoda","score_opus":0.01644484657702985,"score_gpt":0.22774390869411365,"score_spread":0.2112990621170838,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7084253775","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9930379,0.0001924114,0.00018278908,0.000038206956,0.0000042155525,0.000016326481,0.0035320872,0.000017004919,0.0029790278],"genre_scores_gemma":[0.9937995,0.000102532635,0.0003488808,0.000030343062,0.0000012841659,0.000015385021,0.0020805167,0.000011604267,0.0036099802],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9998561,0.000018977866,0.00000627913,0.000051124654,0.00003211521,0.00003540781],"domain_scores_gemma":[0.99928397,0.00022268463,0.00013270149,0.000023283561,0.00021584674,0.00012158516],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00017325117,0.0001128812,0.00024873353,0.0012882898,0.00087995915,0.000463546,0.00038416806,0.00021444236,0.0029669928],"category_scores_gemma":[0.00060934183,0.0001460896,0.0001934421,0.0012727894,0.0003045741,0.0001528285,0.00039481852,0.00014605785,0.00036847618],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036806596,0.000033475495,0.9620072,0.00018031319,0.00014905073,0.00037011807,0.002181552,0.0013277773,0.022368768,0.00023106881,0.0016875113,0.009094967],"study_design_scores_gemma":[0.0000013707912,0.000012306299,0.9982864,0.0000049038354,0.000012964307,0.000025369352,0.0008749113,0.00016458817,0.00020569007,0.000016786582,0.0003917606,0.0000028554869],"about_ca_topic_score_codex":0.7899846,"about_ca_topic_score_gemma":0.94905734,"teacher_disagreement_score":0.21001542,"about_ca_system_score_codex":0.0029559839,"about_ca_system_score_gemma":0.0015127171,"threshold_uncertainty_score":0.42250443},"labels":[],"label_agreement":null},{"id":"W7092316315","doi":"","title":"Schema for In-Context Learning","year":2025,"lang":"","type":"article","venue":"arXiv (Cornell University)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Resources Canada; Canada First Research Excellence Fund; University of Toronto; Natural Sciences and Engineering Research Council of Canada; Government of Canada; Canadian Institute for Advanced Research","keywords":"Schema (genetic algorithms); Cognition; Knowledge representation and reasoning; Abstraction; Mental representation; Transfer of learning; Concept learning; Causal reasoning","score_opus":0.07019211780600276,"score_gpt":0.19527556207749652,"score_spread":0.12508344427149376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7092316315","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035072833,0.00078453356,0.9439663,0.0015923445,0.00014117263,0.0002493567,0.0017311478,0.0076972046,0.00876501],"genre_scores_gemma":[0.41467416,0.00072246976,0.57515544,0.00069685414,0.00007086165,0.00044791796,0.0051423665,0.00046961012,0.0026204141],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99856156,0.0007099289,0.00010412946,0.0003606515,0.00020195842,0.00006175789],"domain_scores_gemma":[0.99611497,0.0013777561,0.00019119267,0.0019227267,0.00025901716,0.0001343652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019013968,0.00072453776,0.0003407042,0.0004972592,0.0003738858,0.0017396276,0.0020454067,0.0008313345,0.0059187426],"category_scores_gemma":[0.010128078,0.0003421662,0.0011450081,0.00061972876,0.0009104312,0.005142901,0.0030736628,0.0023812277,0.0019978439],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030379204,0.00052272,0.0103788385,0.00094007805,0.00019778287,0.00022478479,0.0026255366,0.04374029,0.020113502,0.29015917,0.029019238,0.6017743],"study_design_scores_gemma":[0.00007183698,0.0002066647,0.0015974327,0.00016438341,0.000115417795,0.00040449807,0.00050530036,0.36521,0.024929153,0.48661792,0.12011718,0.000060120517],"about_ca_topic_score_codex":0.0011919031,"about_ca_topic_score_gemma":0.0017430924,"teacher_disagreement_score":0.0059187426,"about_ca_system_score_codex":0.00071554875,"about_ca_system_score_gemma":0.0010738192,"threshold_uncertainty_score":0.019800186},"labels":[],"label_agreement":null},{"id":"W7092316893","doi":"","title":"PRISM: Agentic Retrieval with LLMs for Multi-Hop Question Answering","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Question answering; Context (archaeology); Set (abstract data type); Document retrieval; Language model; Calibration","score_opus":0.06427901365345795,"score_gpt":0.30663004477633915,"score_spread":0.2423510311228812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7092316893","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004075362,0.00090833905,0.962701,0.0007062672,0.00009278463,0.00042618735,0.0010358244,0.026815828,0.0032382607],"genre_scores_gemma":[0.106935106,0.00057594106,0.8791821,0.000854573,0.00014625941,0.00071618025,0.00404287,0.0011448266,0.006402077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973327,0.0012374648,0.00024286626,0.000552727,0.00053842494,0.00009578839],"domain_scores_gemma":[0.9952206,0.0028171798,0.0002310607,0.0010417231,0.00053263677,0.0001566771],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037761668,0.0014983838,0.001097824,0.0021300947,0.0008534026,0.0030253816,0.0030185017,0.0022743088,0.011333261],"category_scores_gemma":[0.01408016,0.00075627136,0.0016832239,0.001378144,0.0010741036,0.006867345,0.0045585493,0.003071933,0.0075234417],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00089178456,0.00067985494,0.0025099532,0.0025448843,0.00040860815,0.0006077123,0.0020901232,0.07226212,0.036545876,0.08866804,0.08656486,0.70622617],"study_design_scores_gemma":[0.00022389308,0.00024636224,0.00042126616,0.0001089398,0.00015860869,0.00033548908,0.00037802142,0.8023853,0.018908577,0.10462648,0.07211018,0.00009685611],"about_ca_topic_score_codex":0.0040124976,"about_ca_topic_score_gemma":0.006613654,"teacher_disagreement_score":0.011333261,"about_ca_system_score_codex":0.0012061783,"about_ca_system_score_gemma":0.0021996736,"threshold_uncertainty_score":0.0379135},"labels":[],"label_agreement":null},{"id":"W7096009128","doi":"","title":"Airline and EpiPen ® Marketer Want Travelers at Risk of Severe Allergic Reactions to Feel Better Prepared","year":2016,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Anaphylactic reactions; Anaphylaxis; Allergic reaction; Risk management","score_opus":0.014127062201752335,"score_gpt":0.22339168644201063,"score_spread":0.2092646242402583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7096009128","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026886288,0.0069656195,0.0028459881,0.6659506,0.026059642,0.0002051753,0.00079697714,0.0018568804,0.26843286],"genre_scores_gemma":[0.091306776,0.0061345827,0.003657266,0.37770277,0.0053957673,0.00010563149,0.0007719727,0.00031273754,0.5146125],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986884,0.0002628482,0.00007320341,0.00014562206,0.00051255105,0.00031744692],"domain_scores_gemma":[0.99722964,0.000332859,0.0003345519,0.00013880154,0.0010376653,0.00092649995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010117639,0.0006480156,0.00027596997,0.0003217726,0.003841432,0.0032046547,0.00048433704,0.004367799,0.10313034],"category_scores_gemma":[0.0051143477,0.00034678218,0.0005974304,0.00026745273,0.001029083,0.0043026595,0.0016702731,0.0040916065,0.021745022],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000073630144,0.00007884203,0.0028896,0.000083903695,0.000013295397,0.0006672741,0.00075232954,0.000024780631,0.00095944345,0.0012109309,0.9601983,0.033047818],"study_design_scores_gemma":[0.00002810954,0.00015833655,0.004127168,0.000100970465,0.00001753924,0.0011832792,0.003967907,0.000081465274,0.0006700599,0.0003813208,0.9892541,0.00002974684],"about_ca_topic_score_codex":0.0142038325,"about_ca_topic_score_gemma":0.022475824,"teacher_disagreement_score":0.10313034,"about_ca_system_score_codex":0.0011095928,"about_ca_system_score_gemma":0.0018774718,"threshold_uncertainty_score":0.34500533},"labels":[],"label_agreement":null},{"id":"W7096448959","doi":"","title":"ton, Canada.","year":2003,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Sentence; Oracle; Compression (physics); Task (project management); Upper and lower bounds","score_opus":0.012550983398717061,"score_gpt":0.19720295760689416,"score_spread":0.18465197420817708,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7096448959","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016732417,0.0060804775,0.004462382,0.007772245,0.0009614897,0.00022482956,0.06513259,0.0019900645,0.8966436],"genre_scores_gemma":[0.03315295,0.002815341,0.0040685963,0.0013775551,0.000048474358,0.00006825817,0.01704228,0.00052028464,0.9409062],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99900573,0.0000717033,0.000035637564,0.00025376317,0.0004216963,0.00021153042],"domain_scores_gemma":[0.99701947,0.0003853808,0.0001311974,0.00024280224,0.0018199474,0.00040125998],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0006998409,0.0011045391,0.00071965676,0.0016416695,0.005370023,0.0048702476,0.0011769661,0.0014317564,0.35976076],"category_scores_gemma":[0.0029485715,0.00050871796,0.00043456568,0.004341405,0.0008760059,0.0014905545,0.0012762181,0.0011850123,0.15560094],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055334327,0.00013344936,0.015181701,0.00071063533,0.000081855025,0.0012010912,0.0011957424,0.0017918238,0.0049137687,0.02612676,0.5764744,0.37163538],"study_design_scores_gemma":[0.000041664982,0.000027814745,0.014730687,0.00015867123,0.000020596963,0.00022212708,0.0012420474,0.0012830695,0.0010013528,0.0021820753,0.9790475,0.000042541076],"about_ca_topic_score_codex":0.83382225,"about_ca_topic_score_gemma":0.9386641,"teacher_disagreement_score":0.64023924,"about_ca_system_score_codex":0.011789348,"about_ca_system_score_gemma":0.017260004,"threshold_uncertainty_score":0.91322356},"labels":[],"label_agreement":null},{"id":"W7096935629","doi":"","title":"Leveraging DUC","year":2006,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Ranking (information retrieval); Scheme (mathematics); Quality (philosophy); Heuristics","score_opus":0.020576338063070007,"score_gpt":0.21894764534265934,"score_spread":0.19837130727958935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7096935629","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057571117,0.006528053,0.7891341,0.0037589662,0.002250161,0.0010517562,0.007872123,0.06473802,0.067095704],"genre_scores_gemma":[0.23649989,0.0020011405,0.702566,0.0007628794,0.0006035805,0.00058142585,0.02564317,0.0039077457,0.02743413],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99211556,0.0040609497,0.0005709224,0.0010981156,0.001901306,0.0002531398],"domain_scores_gemma":[0.98163384,0.0057405047,0.00040772164,0.0046982476,0.007052717,0.00046707148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070758057,0.0013397672,0.0012012275,0.0033954778,0.001373819,0.004130399,0.0015844618,0.0010080708,0.0038860894],"category_scores_gemma":[0.018652644,0.000532364,0.00052379165,0.0022028435,0.00079810526,0.0033291709,0.0019185077,0.00164223,0.003415232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031565575,0.00035685042,0.002080582,0.0007063758,0.00016840159,0.0003249907,0.0012918406,0.040889174,0.02037522,0.015631827,0.11953559,0.79832363],"study_design_scores_gemma":[0.00010524563,0.0005059073,0.0025576372,0.000221216,0.00018211923,0.00052365573,0.00082531193,0.5874799,0.032617643,0.021552674,0.3532231,0.00020566357],"about_ca_topic_score_codex":0.02881272,"about_ca_topic_score_gemma":0.05985315,"teacher_disagreement_score":0.02881272,"about_ca_system_score_codex":0.0017429721,"about_ca_system_score_gemma":0.0024537502,"threshold_uncertainty_score":0.057290018},"labels":[],"label_agreement":null},{"id":"W7097187097","doi":"","title":"Simple and Effective Question Processing using Regular Expressions and WordNet","year":2005,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Question answering; WordNet; Simplicity; Component (thermodynamics); Simple (philosophy); Word (group theory); Debugging","score_opus":0.015289628275752074,"score_gpt":0.2836651031731034,"score_spread":0.2683754748973513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7097187097","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008087631,0.00019387911,0.97056085,0.0004768373,0.000042625546,0.0002869816,0.0006884725,0.01598822,0.0036744578],"genre_scores_gemma":[0.093198426,0.00040486673,0.8951518,0.00027626727,0.00011473823,0.00029808038,0.003542766,0.0018390964,0.0051739262],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99589586,0.0015018967,0.000472485,0.000952007,0.0009910556,0.0001866425],"domain_scores_gemma":[0.9928276,0.0037871904,0.00048053055,0.0016911546,0.0010632642,0.0001502084],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0041687754,0.0012586046,0.0013977248,0.0031701159,0.0010937717,0.0036113828,0.0020509972,0.0012815008,0.0063657146],"category_scores_gemma":[0.0121096,0.0010975118,0.0018478974,0.002338452,0.0013525815,0.010867243,0.0026325139,0.0018041468,0.0065263994],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038789093,0.00040188717,0.0044895536,0.0009963878,0.00020887953,0.00051985896,0.0035113078,0.017626155,0.055044744,0.16609843,0.028628994,0.7220861],"study_design_scores_gemma":[0.00010097927,0.00018714281,0.002367202,0.0002129196,0.00018232799,0.0008839107,0.0011412912,0.44311413,0.07140854,0.34350666,0.13671441,0.0001804215],"about_ca_topic_score_codex":0.0026429081,"about_ca_topic_score_gemma":0.0034729869,"teacher_disagreement_score":0.0063657146,"about_ca_system_score_codex":0.00090709364,"about_ca_system_score_gemma":0.0014861659,"threshold_uncertainty_score":0.022046864},"labels":[],"label_agreement":null},{"id":"W7098422109","doi":"","title":"Preliminary and Incomplete.","year":2002,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Autoregressive model; Estimator; Volatility (finance); Estimation; Scaling; Process (computing); Econometric model; Series (stratigraphy)","score_opus":0.04153136491727264,"score_gpt":0.22308153692069252,"score_spread":0.18155017200341989,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098422109","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025867207,0.011471707,0.01333238,0.020663168,0.01258734,0.0005071752,0.11421966,0.0022905427,0.82234126],"genre_scores_gemma":[0.020641483,0.008553256,0.006029868,0.0033568535,0.0020155013,0.00057379395,0.04875107,0.0010103711,0.9090679],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983895,0.00039892158,0.0001837787,0.0003548532,0.00052902906,0.00014398951],"domain_scores_gemma":[0.9959502,0.0009714423,0.00026284045,0.0009942949,0.0015385416,0.00028253224],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.002091559,0.0013516917,0.0011762382,0.003175791,0.0019034857,0.0038879043,0.0017051372,0.0017086879,0.65188193],"category_scores_gemma":[0.010738111,0.0006595882,0.00085780624,0.005171672,0.00079824065,0.0049517625,0.0033413663,0.0017548316,0.37586933],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000078898265,0.000023730754,0.00039669062,0.00030231962,0.000013638385,0.00008083,0.00020406226,0.0001427285,0.00013443157,0.026753454,0.9006841,0.07118505],"study_design_scores_gemma":[0.000008361057,0.000007795115,0.00047807765,0.00010930346,0.0000036741753,0.000056442976,0.00009939266,0.0001003025,0.00006580333,0.004257473,0.9948078,0.0000056987014],"about_ca_topic_score_codex":0.008069087,"about_ca_topic_score_gemma":0.010877561,"teacher_disagreement_score":0.34811807,"about_ca_system_score_codex":0.0029683148,"about_ca_system_score_gemma":0.0027430241,"threshold_uncertainty_score":0.49654818},"labels":[],"label_agreement":null},{"id":"W7099458630","doi":"","title":"Submitted by:","year":2002,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.027297456615216962,"score_gpt":0.20597463197171823,"score_spread":0.17867717535650127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099458630","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010478827,0.004257747,0.0015899108,0.059370343,0.05531946,0.0010193986,0.012026083,0.0013684612,0.85456973],"genre_scores_gemma":[0.004751263,0.00043119196,0.00020553477,0.0009293522,0.00075925514,0.000022517406,0.00072894676,0.00013784616,0.992034],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99940884,0.000042139258,0.000041890275,0.00016707665,0.00021774034,0.00012227647],"domain_scores_gemma":[0.9961553,0.00045036242,0.00006568134,0.00017594044,0.002362689,0.00078998634],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00084036466,0.0006587207,0.0008679963,0.0008093879,0.0017995018,0.004091585,0.0009781506,0.0030023453,0.71497846],"category_scores_gemma":[0.0057986686,0.00029769808,0.0004666065,0.00052322634,0.0008417535,0.0010752497,0.0011279393,0.0017776798,0.37467763],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022179385,0.000029695959,0.0004142837,0.00013871591,0.0000099933195,0.00015430007,0.00009277293,0.00010269429,0.0005850772,0.0012840986,0.9745832,0.02238333],"study_design_scores_gemma":[0.000024272576,0.000035794485,0.0012712263,0.000086005995,0.000006837833,0.000048256727,0.00016951544,0.0001391106,0.0002778321,0.00032649975,0.9976083,0.000006405418],"about_ca_topic_score_codex":0.017856555,"about_ca_topic_score_gemma":0.060681783,"teacher_disagreement_score":0.28502154,"about_ca_system_score_codex":0.0022728473,"about_ca_system_score_gemma":0.0020184154,"threshold_uncertainty_score":0.40654862},"labels":[],"label_agreement":null},{"id":"W7099778194","doi":"","title":"Are Government Spending Multipliers Greater During Times of Slack? Evidence from 20th Century Historical Data.” American Economic Review 103(2","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Newspaper; Government (linguistics); Subject (documents); Government spending; Developed country; Economic forecasting","score_opus":0.052778178422459764,"score_gpt":0.25943654786962883,"score_spread":0.20665836944716906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099778194","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5059896,0.21559602,0.0016354873,0.13263392,0.0017936602,0.000060427305,0.03956508,0.00010761479,0.10261835],"genre_scores_gemma":[0.9392014,0.041952867,0.00035555472,0.003283322,0.001250737,0.000028755796,0.008585371,0.000054108645,0.005287912],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99905866,0.0003347367,0.00007917294,0.00014101833,0.00020912515,0.00017729697],"domain_scores_gemma":[0.97565633,0.007500925,0.010824715,0.0012361764,0.0040222877,0.0007596244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043014255,0.0002804959,0.0005067806,0.0037668417,0.0006497642,0.002299481,0.0007583433,0.0011518159,0.007820624],"category_scores_gemma":[0.021102408,0.00037163566,0.0005419217,0.0074068382,0.0010627644,0.0021654076,0.0008261698,0.0013261164,0.0014197681],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012403494,0.0001156484,0.602202,0.0011589881,0.0016405013,0.00056022,0.002102287,0.0020519372,0.0002129795,0.043949887,0.22271892,0.12204628],"study_design_scores_gemma":[0.00013692037,0.000066599074,0.84543633,0.00086409884,0.0007407396,0.00015056014,0.0016937581,0.0013849967,0.00055299705,0.008519322,0.14039408,0.000059666694],"about_ca_topic_score_codex":0.06585569,"about_ca_topic_score_gemma":0.0932311,"teacher_disagreement_score":0.06585569,"about_ca_system_score_codex":0.0019187279,"about_ca_system_score_gemma":0.0019648438,"threshold_uncertainty_score":0.13094473},"labels":[],"label_agreement":null},{"id":"W7099875465","doi":"","title":"ACCEPTED BY IEEE TRANSACTIONS ON POWER SYSTEMS 1","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Cascading failure; Electric power system; Power-system protection; Generator (circuit theory); Sensitivity (control systems); Ambiguity; Power (physics); AC power; Key (lock)","score_opus":0.016384974559029254,"score_gpt":0.2205543175482751,"score_spread":0.20416934298924583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099875465","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0063373605,0.012428377,0.16544586,0.014356992,0.036262248,0.0003528663,0.003399139,0.0033255606,0.75809157],"genre_scores_gemma":[0.13421038,0.011356054,0.029124392,0.002183746,0.0060463357,0.00027832272,0.005882887,0.0010520051,0.8098659],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990657,0.00019164951,0.00009045106,0.00021916574,0.00035720976,0.000075950695],"domain_scores_gemma":[0.99798334,0.000271155,0.00011166908,0.00046037257,0.0010380683,0.00013526285],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011466269,0.0010734008,0.0010407859,0.0009229673,0.00088975276,0.004149161,0.0012444415,0.0017620005,0.24475452],"category_scores_gemma":[0.0036098696,0.0003762304,0.00080127304,0.0021645196,0.0006819788,0.002446522,0.001755901,0.0019996853,0.09328479],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013258398,0.00009203349,0.0012429329,0.00047269044,0.00008440371,0.00032476653,0.00017819546,0.009581293,0.0016632763,0.071015045,0.511428,0.40378478],"study_design_scores_gemma":[0.00002722671,0.0000686618,0.00071605865,0.00017428143,0.000025090696,0.0002234898,0.00010268709,0.012369028,0.0004644399,0.024380852,0.9614306,0.000017482518],"about_ca_topic_score_codex":0.0022989921,"about_ca_topic_score_gemma":0.0023404842,"teacher_disagreement_score":0.75524545,"about_ca_system_score_codex":0.00083649915,"about_ca_system_score_gemma":0.0013094329,"threshold_uncertainty_score":0.8187854},"labels":[],"label_agreement":null},{"id":"W7100283491","doi":"","title":": Canada (2013)&amp;quot; ROBUST TREE-STRUCTURED NAMED ENTITIES RECOGNITION FROM SPEECH","year":2013,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Named-entity recognition; Entity linking; Conditional random field; Set (abstract data type); Knowledge base; Conjunction (astronomy); Question answering; Tree (set theory)","score_opus":0.029847109839985876,"score_gpt":0.1954199673362026,"score_spread":0.16557285749621672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7100283491","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027230533,0.013252823,0.052858368,0.07828835,0.0133171175,0.00051263435,0.20202005,0.013057592,0.59946245],"genre_scores_gemma":[0.06085345,0.0033096215,0.022982692,0.0036411046,0.0003927015,0.000066152075,0.04054809,0.0013246764,0.8668816],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994216,0.000027463468,0.000014264256,0.000101004276,0.0003345043,0.00010114627],"domain_scores_gemma":[0.99884665,0.00006241007,0.000028272863,0.00006854609,0.0008427016,0.00015151677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012221605,0.0005107918,0.0003764277,0.0013888705,0.0030414115,0.0036480126,0.00090444164,0.0014071581,0.054069184],"category_scores_gemma":[0.0013422052,0.00027801612,0.00040215443,0.0022105933,0.0009499255,0.0011509313,0.0006862859,0.0010247764,0.018393125],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011683356,0.000027987528,0.0022981004,0.000119233715,0.000018654404,0.00024224135,0.00018286347,0.0013238105,0.0030955598,0.008582932,0.8903119,0.093679845],"study_design_scores_gemma":[0.000016155638,0.000012186544,0.006950883,0.000044136763,0.000009354724,0.000061127415,0.00023503142,0.0023198319,0.0042606816,0.0019155595,0.98413587,0.00003921302],"about_ca_topic_score_codex":0.8752568,"about_ca_topic_score_gemma":0.9458094,"teacher_disagreement_score":0.12474322,"about_ca_system_score_codex":0.014223897,"about_ca_system_score_gemma":0.019224403,"threshold_uncertainty_score":0.2509557},"labels":[],"label_agreement":null},{"id":"W7102328961","doi":"","title":"Understanding In-Context Learning Beyond Transformers: An Investigation of State Space and Hybrid Architectures","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada","keywords":"Task (project management); Function (biology); Space (punctuation); State (computer science); Hybrid learning; Mechanism (biology)","score_opus":0.06771266374553078,"score_gpt":0.26477068188613023,"score_spread":0.19705801814059945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7102328961","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6262074,0.0014771532,0.3580006,0.001476861,0.000107285465,0.00020493673,0.00041145788,0.0038127394,0.008301558],"genre_scores_gemma":[0.9573182,0.00020196498,0.04072161,0.00017231148,0.000017663398,0.00008754,0.0003054109,0.00012170439,0.0010535872],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982613,0.0010521836,0.000068994574,0.00031468124,0.00017655401,0.0001262846],"domain_scores_gemma":[0.9899659,0.0075125005,0.00028843363,0.0014865631,0.00048456815,0.0002620294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005017604,0.0010565983,0.00074249937,0.0004548862,0.00034757884,0.0020516445,0.001672411,0.0010676638,0.0028974365],"category_scores_gemma":[0.018512024,0.00029733556,0.00075850845,0.00040052342,0.0010858156,0.0062833005,0.0020514445,0.0036917513,0.0006102508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013473874,0.0010518834,0.015619544,0.00095758255,0.000384682,0.00020848005,0.0018304258,0.53854966,0.018448329,0.03894473,0.0038626387,0.3787948],"study_design_scores_gemma":[0.000028169261,0.0002809303,0.00056864443,0.000019394636,0.000039989387,0.000025695044,0.00014777509,0.9710462,0.0035630956,0.023321897,0.0009434198,0.000014834576],"about_ca_topic_score_codex":0.004645875,"about_ca_topic_score_gemma":0.006167086,"teacher_disagreement_score":0.005017604,"about_ca_system_score_codex":0.001035104,"about_ca_system_score_gemma":0.0010510966,"threshold_uncertainty_score":0.026535988},"labels":[],"label_agreement":null},{"id":"W7103204130","doi":"","title":"LoRAQuant: Mixed-Precision Quantization of LoRA to Ultra-Low Bits","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Mitacs","keywords":"Quantization (signal processing); Adapter (computing); Row; Automatic summarization; Vector quantization; Personalization; Naturalness","score_opus":0.04060280214240603,"score_gpt":0.29189316236988433,"score_spread":0.2512903602274783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7103204130","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012339758,0.001067797,0.9715484,0.00051047193,0.00027793954,0.00012321713,0.00058730954,0.011342221,0.0022029371],"genre_scores_gemma":[0.28863385,0.0006637789,0.6978954,0.000967637,0.0002305198,0.00044825382,0.0022795196,0.0010808868,0.0078001176],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986437,0.0004307714,0.000100167286,0.00028601917,0.00043535797,0.00010408381],"domain_scores_gemma":[0.99740726,0.0010180506,0.0001485234,0.0007242868,0.0006041606,0.000097708435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018642305,0.0013780259,0.00093059405,0.00091559195,0.00062540005,0.0015776099,0.0017378209,0.0010711168,0.008505591],"category_scores_gemma":[0.010880855,0.00047037692,0.000693206,0.0009781297,0.0008646301,0.0026499827,0.0021886807,0.0026171408,0.004121887],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007990514,0.00023103942,0.0012329607,0.0003825489,0.0001245499,0.00014754236,0.00033017204,0.11351161,0.028402783,0.024210392,0.03815506,0.7924724],"study_design_scores_gemma":[0.000087026805,0.00016526271,0.0003800087,0.00004181206,0.000021066031,0.00009707357,0.000075865675,0.9573654,0.013301554,0.019957991,0.008462201,0.00004477381],"about_ca_topic_score_codex":0.0050472766,"about_ca_topic_score_gemma":0.011355031,"teacher_disagreement_score":0.008505591,"about_ca_system_score_codex":0.00079777377,"about_ca_system_score_gemma":0.0012355967,"threshold_uncertainty_score":0.028454006},"labels":[],"label_agreement":null},{"id":"W7104370506","doi":"10.18653/v1/2025.newsum-main.2","title":"Hierarchical Attention Adapter for Abstractive Dialogue Summarization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Adapter (computing); Key (lock); Redundancy (engineering)","score_opus":0.027473027713074534,"score_gpt":0.28005662229897194,"score_spread":0.2525835945858974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7104370506","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03672819,0.0023489299,0.89077365,0.0006181401,0.00028944356,0.00034510012,0.0017676551,0.061414327,0.005714608],"genre_scores_gemma":[0.51608926,0.0009283892,0.44951865,0.0011263284,0.00044871776,0.00085349154,0.011959247,0.0019166454,0.017159296],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99876803,0.00047588308,0.000078273006,0.000371722,0.00018981754,0.000116349824],"domain_scores_gemma":[0.99868315,0.000559253,0.00008358589,0.00034267927,0.00024270137,0.00008860118],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016406485,0.001646443,0.0011260755,0.001517778,0.000656969,0.0011237449,0.0020604392,0.0013643288,0.0072161495],"category_scores_gemma":[0.005067232,0.00038544743,0.0010500231,0.0011120935,0.00048772947,0.0026574642,0.0022733212,0.0022010354,0.004509491],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062106375,0.00026890123,0.001280254,0.0006147191,0.0001759177,0.00017380124,0.00092238665,0.03643175,0.028913533,0.0070774835,0.030397259,0.8931229],"study_design_scores_gemma":[0.00015561904,0.0005460984,0.0019364378,0.000076405064,0.00022946838,0.00020433488,0.0005944259,0.8964337,0.035465203,0.022691771,0.04159026,0.000076250166],"about_ca_topic_score_codex":0.006271385,"about_ca_topic_score_gemma":0.010574818,"teacher_disagreement_score":0.0072161495,"about_ca_system_score_codex":0.0010754254,"about_ca_system_score_gemma":0.0012265746,"threshold_uncertainty_score":0.024140418},"labels":[],"label_agreement":null},{"id":"W7104481854","doi":"10.71781/10885","title":"Towards efficient large language models : training low-bitwidth variants and low-rank decomposition of pretrained models","year":2024,"lang":"en","type":"dissertation","venue":"Open MIND","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Oak Ridge National Laboratory; Office of Science; Canadian Institute for Advanced Research; Canada Excellence Research Chairs, Government of Canada; U.S. Department of Energy","keywords":"Training set; Statistical analysis; Context (archaeology)","score_opus":0.03243441090778062,"score_gpt":0.3186131654673854,"score_spread":0.2861787545596048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7104481854","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055801496,0.0013897902,0.9338637,0.00053255656,0.00016717282,0.00007827466,0.00047682915,0.0056865844,0.0020036506],"genre_scores_gemma":[0.49567387,0.001070469,0.48544097,0.00077171775,0.00021886738,0.0003686253,0.0041455794,0.0011992949,0.011110611],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99924636,0.00023905994,0.00004732712,0.00021806297,0.0001365275,0.000112761736],"domain_scores_gemma":[0.99788755,0.0011849522,0.00012718943,0.00037458178,0.00032171182,0.00010410563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013552686,0.0021512848,0.001452073,0.0008132062,0.00039499413,0.0018263246,0.002000582,0.0017459225,0.003962906],"category_scores_gemma":[0.006398468,0.00094020675,0.0016862811,0.0008652358,0.00063044834,0.0028674982,0.0013077466,0.003946149,0.0029310626],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004456111,0.00025009332,0.0013251398,0.00022694955,0.00019886774,0.0001812567,0.00019350769,0.6283007,0.015097006,0.0059096236,0.0071076765,0.34076354],"study_design_scores_gemma":[0.0000090644535,0.000035478955,0.000089536006,0.0000091265,0.000010660824,0.000015036524,0.000015964208,0.9962998,0.0012350827,0.0018743522,0.00040013433,0.0000057118764],"about_ca_topic_score_codex":0.011679534,"about_ca_topic_score_gemma":0.018665181,"teacher_disagreement_score":0.011679534,"about_ca_system_score_codex":0.00085910526,"about_ca_system_score_gemma":0.0013857854,"threshold_uncertainty_score":0.023223102},"labels":[],"label_agreement":null},{"id":"W7106288743","doi":"10.36227/techrxiv.176369738.87142789/v1","title":"ContextMentalQA: Modeling Cultural, Social, and Religious Context in Arabic Mental Health Question Answering","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Mental health; Context (archaeology); Distress; Arabic; Question answering; Mental distress","score_opus":0.036324455342996845,"score_gpt":0.3288648955707391,"score_spread":0.2925404402277423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106288743","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39199942,0.004023755,0.503812,0.006888457,0.0008793439,0.0025433425,0.056692064,0.019823851,0.0133377],"genre_scores_gemma":[0.5852305,0.00064334564,0.32738188,0.0012448676,0.00022799439,0.0019225336,0.0768905,0.0004101246,0.006048341],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837995,0.0007862809,0.00011137349,0.0004937149,0.00013943622,0.00008915038],"domain_scores_gemma":[0.99692553,0.0019770171,0.00014469365,0.0003511246,0.00046853663,0.0001331042],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002534054,0.0012258914,0.00046256723,0.0016540092,0.0010803024,0.0017193623,0.0014874012,0.0015472317,0.0040217675],"category_scores_gemma":[0.008336054,0.00030944357,0.0010752897,0.0010325989,0.0006507309,0.0024248126,0.0026233378,0.0021680477,0.002066234],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016280374,0.0011919021,0.062431205,0.0022154718,0.00031283387,0.0009100303,0.010689034,0.104565896,0.03054905,0.023802275,0.12113476,0.6405695],"study_design_scores_gemma":[0.000121959296,0.00022785386,0.02180456,0.00023765174,0.00011464295,0.00043864103,0.0034892375,0.8556387,0.0132601,0.026581788,0.07798169,0.000103266415],"about_ca_topic_score_codex":0.021154512,"about_ca_topic_score_gemma":0.034748312,"teacher_disagreement_score":0.021154512,"about_ca_system_score_codex":0.0018432777,"about_ca_system_score_gemma":0.0015082109,"threshold_uncertainty_score":0.04206276},"labels":[],"label_agreement":null},{"id":"W7106288924","doi":"10.2139/ssrn.5763184","title":"Efficient Inference Using Large Language Models with Limited Human Data: Fine-Tuning then Rectification","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Quest University Canada","funders":"","keywords":"Rectification; Inference; Language model; Data modeling; Natural language","score_opus":0.06451865640085883,"score_gpt":0.32517438511996394,"score_spread":0.26065572871910514,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106288924","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015242392,0.0008519496,0.97745854,0.0011264707,0.00012144727,0.000075851836,0.00037315863,0.00399532,0.00075495074],"genre_scores_gemma":[0.43782112,0.0007983557,0.54893357,0.0011493806,0.0006599762,0.00039110892,0.0029257722,0.0014007979,0.0059199785],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9962794,0.0015706675,0.0002529635,0.0012735912,0.0003204233,0.00030291133],"domain_scores_gemma":[0.9820539,0.013592292,0.0004691673,0.0028877517,0.0007178735,0.0002789658],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0083811665,0.0014452402,0.0029680072,0.0015367899,0.0012388728,0.002973716,0.004103714,0.0035120558,0.0072902367],"category_scores_gemma":[0.038609017,0.0019907705,0.0021127458,0.0019493771,0.0015119079,0.0055694575,0.0038415038,0.0064014946,0.004239972],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015060644,0.00036144632,0.0048346873,0.000586348,0.00051075383,0.00035621,0.00081090524,0.17940667,0.016328298,0.019010687,0.016801713,0.7594862],"study_design_scores_gemma":[0.00009042906,0.00004708242,0.000728043,0.000027473723,0.00007923234,0.00009221561,0.00008800995,0.95097053,0.0030176863,0.042700183,0.002131118,0.000028096934],"about_ca_topic_score_codex":0.011929025,"about_ca_topic_score_gemma":0.017502472,"teacher_disagreement_score":0.011929025,"about_ca_system_score_codex":0.0010228207,"about_ca_system_score_gemma":0.0030183138,"threshold_uncertainty_score":0.044324398},"labels":[],"label_agreement":null},{"id":"W7106480312","doi":"10.1609/aaaiss.v7i1.36923","title":"Fine-Tuning Large Language Models for Structured ClinicalReport Generation Using GRPO","year":2025,"lang":"","type":"article","venue":"Proceedings of the AAAI Symposium Series","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"International Institute for Sustainable Development; Western University","funders":"","keywords":"Relevance (law); Adaptation (eye); Language model; Disk formatting; Medical care; Baseline (sea); Complement (music); English language","score_opus":0.03647737372343484,"score_gpt":0.29663727927137035,"score_spread":0.2601599055479355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106480312","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14695507,0.0030078548,0.77581704,0.0040996624,0.00079848565,0.0007548276,0.0044055143,0.05724932,0.0069122063],"genre_scores_gemma":[0.6836118,0.00059255137,0.29981688,0.0019861362,0.00019766948,0.0006623388,0.0074713,0.0012545318,0.0044067977],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99810386,0.0008628741,0.000120907105,0.0005358977,0.00024159903,0.00013481051],"domain_scores_gemma":[0.99240667,0.00536517,0.00034487178,0.0008912735,0.0006969343,0.0002951588],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004664051,0.0016616494,0.0009829941,0.0008324966,0.0003392246,0.0014172595,0.0023148993,0.0016921584,0.003183582],"category_scores_gemma":[0.019038444,0.00066712947,0.0010976718,0.0007020774,0.00069215405,0.0022355833,0.0015000077,0.0032632933,0.0023666092],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000722271,0.00064729725,0.0058634044,0.00043292597,0.00021381535,0.00035047447,0.00028409882,0.6532741,0.0080912355,0.004141379,0.029242104,0.29673687],"study_design_scores_gemma":[0.000085690386,0.00007537634,0.00019705809,0.000016864831,0.000021923925,0.000043024065,0.000032489184,0.9909603,0.0035165907,0.0031917917,0.0018432032,0.000015756732],"about_ca_topic_score_codex":0.008011099,"about_ca_topic_score_gemma":0.013609675,"teacher_disagreement_score":0.008011099,"about_ca_system_score_codex":0.0016114939,"about_ca_system_score_gemma":0.003240456,"threshold_uncertainty_score":0.02466613},"labels":[],"label_agreement":null},{"id":"W7106484130","doi":"10.1609/aaaiss.v7i1.36936","title":"Hermes: A Modular Multi-Agent System for StructuringClinical Text","year":2025,"lang":"","type":"article","venue":"Proceedings of the AAAI Symposium Series","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Modular design; Unstructured data; Structuring; Bridging (networking); Graph; Architecture; Knowledge graph; Natural language; Semantic network","score_opus":0.020438566578698265,"score_gpt":0.26123595357397716,"score_spread":0.2407973869952789,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106484130","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007902598,0.00014243476,0.9518056,0.00074091414,0.00014039579,0.0010954184,0.0008590972,0.03178314,0.005530338],"genre_scores_gemma":[0.08739998,0.000244725,0.89756787,0.00041498162,0.00008172681,0.0011628452,0.0022688212,0.0010634603,0.0097956965],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99848336,0.00063726277,0.00017708895,0.00027123626,0.00037057098,0.00006050403],"domain_scores_gemma":[0.9960576,0.002413783,0.0002726886,0.0005878368,0.0004209514,0.00024705706],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039536078,0.0008183706,0.00044970406,0.0011213832,0.00081446796,0.0019481935,0.0020412246,0.0011580532,0.007889927],"category_scores_gemma":[0.010348692,0.0005672676,0.00078801427,0.00047839983,0.0007883707,0.0024398086,0.0031084085,0.0012549231,0.002601553],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014428291,0.00093926943,0.005339884,0.0010704112,0.00033636467,0.0014439339,0.004304055,0.109516926,0.03397217,0.10730834,0.080075465,0.65425044],"study_design_scores_gemma":[0.00047796694,0.00030171112,0.0011720465,0.000168382,0.00014050782,0.00043403538,0.0004890576,0.62874186,0.021777358,0.083679,0.26246607,0.00015199489],"about_ca_topic_score_codex":0.0030723773,"about_ca_topic_score_gemma":0.003721666,"teacher_disagreement_score":0.007889927,"about_ca_system_score_codex":0.00077439117,"about_ca_system_score_gemma":0.002312522,"threshold_uncertainty_score":0.026394427},"labels":[],"label_agreement":null},{"id":"W7108315141","doi":"10.1145/3767695.3769494","title":"FalseCoTQA: Adversarial Multi-Hop QA via Knowledge-Grounded False Chains of Thought","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Guelph","funders":"","keywords":"Adversarial system; Inference; Robustness (evolution); Construct (python library); Benchmark (surveying); Knowledge graph; Question answering; Language model","score_opus":0.03161552785915398,"score_gpt":0.29671721564893705,"score_spread":0.26510168778978305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7108315141","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09251154,0.0021009797,0.8627113,0.0048923087,0.0005388151,0.00058762677,0.003925903,0.018656114,0.014075323],"genre_scores_gemma":[0.7537865,0.00039833883,0.23179881,0.0019289505,0.00018705828,0.00035506027,0.004460916,0.00084260653,0.006241729],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99421906,0.0029515536,0.00023164698,0.0010030287,0.0012693559,0.00032538117],"domain_scores_gemma":[0.97622114,0.017058875,0.0007952292,0.004250195,0.0012354197,0.00043916728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005278556,0.0012002765,0.0007968522,0.0008247292,0.0010469182,0.0021774087,0.003399818,0.002865327,0.0058632037],"category_scores_gemma":[0.033603672,0.00046666362,0.0011487773,0.0006114842,0.0023723387,0.004795083,0.0044895546,0.0038304506,0.0017176919],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013316399,0.0006302282,0.008014038,0.0014796358,0.00041281641,0.0008221477,0.0016447164,0.5866715,0.020478982,0.14544553,0.06727143,0.16579732],"study_design_scores_gemma":[0.0000722754,0.00012480306,0.00030175352,0.00004616555,0.000028654127,0.00015270089,0.00011110339,0.9219092,0.0050465213,0.064463854,0.0077117872,0.000031087573],"about_ca_topic_score_codex":0.005364001,"about_ca_topic_score_gemma":0.007325039,"teacher_disagreement_score":0.0058632037,"about_ca_system_score_codex":0.0015531013,"about_ca_system_score_gemma":0.0020015142,"threshold_uncertainty_score":0.027916014},"labels":[],"label_agreement":null},{"id":"W7110016663","doi":"10.21428/594757db.13f92bbd","title":"Let’s CONFER: A Dataset for Evaluating Natural Language Inference Models on CONditional InFERence and Presupposition","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Carleton University","funders":"","keywords":"Presupposition; Inference; Generalization; Task (project management); Sentence; Natural language","score_opus":0.06303050850121361,"score_gpt":0.3912276692429181,"score_spread":0.32819716074170446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7110016663","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10113228,0.0042312755,0.038875103,0.0026192446,0.0010502809,0.0014254357,0.7705444,0.056977585,0.023144403],"genre_scores_gemma":[0.079348035,0.00043112237,0.044005636,0.00058119063,0.00013642268,0.0006449248,0.86952394,0.0006586557,0.0046701166],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99786973,0.000606296,0.00020615476,0.00068939925,0.0004983451,0.00013009956],"domain_scores_gemma":[0.9942821,0.002824751,0.00031591946,0.0013417457,0.00082601,0.0004095355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002382319,0.0025460022,0.00084825617,0.0023905684,0.0015645237,0.0018714009,0.0040955255,0.003531096,0.010528907],"category_scores_gemma":[0.015527408,0.0005157691,0.0017962902,0.0021195456,0.0009901843,0.0036348691,0.002432112,0.0034455964,0.0070182728],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016382253,0.001566895,0.013391836,0.003394244,0.00056625216,0.0008423409,0.00073191937,0.021148328,0.008017856,0.009416594,0.8209827,0.11830282],"study_design_scores_gemma":[0.0018415409,0.0010810725,0.04291297,0.00061696605,0.00048539683,0.001867583,0.0016476105,0.33001384,0.02838642,0.030474737,0.56023544,0.00043652582],"about_ca_topic_score_codex":0.023986531,"about_ca_topic_score_gemma":0.062086858,"teacher_disagreement_score":0.023986531,"about_ca_system_score_codex":0.0019624524,"about_ca_system_score_gemma":0.0023520407,"threshold_uncertainty_score":0.04769379},"labels":[],"label_agreement":null},{"id":"W7110836110","doi":"10.1080/10705511.2025.2588572","title":"Evaluating Approaches for the Handling of Sign Reflection in Bayesian Latent Variable Models","year":2025,"lang":"en","type":"article","venue":"Structural Equation Modeling A Multidisciplinary Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Reflection (computer programming); Latent variable; Bayesian probability; Variable (mathematics); Sign (mathematics); Pattern recognition (psychology)","score_opus":0.20543495730840183,"score_gpt":0.37482561792634794,"score_spread":0.1693906606179461,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7110836110","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020514477,0.00091809384,0.9748574,0.0009717996,0.00008982398,0.00029437465,0.00012498375,0.0006790401,0.0015499964],"genre_scores_gemma":[0.13889438,0.0005291473,0.8579695,0.00037987853,0.00008526502,0.00087351113,0.0003812708,0.00047621143,0.00041088182],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.91338617,0.07361551,0.0032696745,0.0038059838,0.0050245165,0.00089811074],"domain_scores_gemma":[0.4218009,0.53896886,0.0098289475,0.017303856,0.009921358,0.002176136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.15940392,0.0030422849,0.0021352994,0.005484968,0.0024995285,0.006081527,0.0055020656,0.005543518,0.0063675484],"category_scores_gemma":[0.5267564,0.0023841418,0.0032891768,0.0047364947,0.0038482994,0.010359197,0.0073969583,0.0074332156,0.001146439],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011711563,0.0005031063,0.028956443,0.0015955442,0.0022697852,0.00032939197,0.0031504598,0.373786,0.0015019729,0.21415986,0.0062630954,0.36631316],"study_design_scores_gemma":[0.00025748924,0.00027794894,0.0018735824,0.00059484714,0.00026344918,0.00016702975,0.0004742279,0.79962665,0.0014435183,0.1920176,0.0028597408,0.00014386485],"about_ca_topic_score_codex":0.008095411,"about_ca_topic_score_gemma":0.011312181,"teacher_disagreement_score":0.15940392,"about_ca_system_score_codex":0.0038654432,"about_ca_system_score_gemma":0.0058375387,"threshold_uncertainty_score":0.8430186},"labels":[],"label_agreement":null},{"id":"W7110925557","doi":"10.5281/zenodo.17822780","title":"Long-Text Summarisation","year":2025,"lang":"","type":"dissertation","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Northern Alberta Institute of Technology","funders":"","keywords":"Automatic summarization; Coreference; Transformer; Bridging (networking); Headline","score_opus":0.042884284256458426,"score_gpt":0.26765216900327704,"score_spread":0.2247678847468186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7110925557","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035196956,0.0058273673,0.9155658,0.00085407065,0.00089467684,0.000402476,0.0060517793,0.029247481,0.0059593874],"genre_scores_gemma":[0.33704183,0.003083557,0.5957798,0.00048318066,0.0010132799,0.00043685842,0.03965754,0.0025223554,0.019981647],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99897194,0.0002826898,0.00009986373,0.00035875768,0.00021231813,0.000074366806],"domain_scores_gemma":[0.99675816,0.0013454219,0.000285176,0.0005894859,0.00092813076,0.00009365872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016798219,0.001993202,0.0010120561,0.0028960705,0.0005158682,0.0018046303,0.0014940802,0.000928265,0.006526666],"category_scores_gemma":[0.005768757,0.0003548353,0.0012857359,0.002347688,0.0002510058,0.0029738958,0.0010932877,0.001231864,0.0063375067],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005014084,0.00015511666,0.0014521856,0.0011158264,0.00032213522,0.00025277282,0.00039361534,0.035379965,0.050165046,0.004707364,0.035739116,0.86981547],"study_design_scores_gemma":[0.00016308484,0.0009709016,0.0048815836,0.00017642944,0.00072936295,0.000672077,0.0005392967,0.7267625,0.12168091,0.022620382,0.12067318,0.00013036413],"about_ca_topic_score_codex":0.0019375908,"about_ca_topic_score_gemma":0.0043482883,"teacher_disagreement_score":0.006526666,"about_ca_system_score_codex":0.0005501505,"about_ca_system_score_gemma":0.0007433454,"threshold_uncertainty_score":0.021833897},"labels":[],"label_agreement":null},{"id":"W7111014860","doi":"10.5281/zenodo.17822781","title":"Long-Text Summarisation","year":2025,"lang":"","type":"dissertation","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Northern Alberta Institute of Technology","funders":"","keywords":"Automatic summarization; Coreference; Transformer; Bridging (networking); Headline","score_opus":0.042884284256458426,"score_gpt":0.26765216900327704,"score_spread":0.2247678847468186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7111014860","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035196956,0.0058273673,0.9155658,0.00085407065,0.00089467684,0.000402476,0.0060517793,0.029247481,0.0059593874],"genre_scores_gemma":[0.33704183,0.003083557,0.5957798,0.00048318066,0.0010132799,0.00043685842,0.03965754,0.0025223554,0.019981647],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99897194,0.0002826898,0.00009986373,0.00035875768,0.00021231813,0.000074366806],"domain_scores_gemma":[0.99675816,0.0013454219,0.000285176,0.0005894859,0.00092813076,0.00009365872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016798219,0.001993202,0.0010120561,0.0028960705,0.0005158682,0.0018046303,0.0014940802,0.000928265,0.006526666],"category_scores_gemma":[0.005768757,0.0003548353,0.0012857359,0.002347688,0.0002510058,0.0029738958,0.0010932877,0.001231864,0.0063375067],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005014084,0.00015511666,0.0014521856,0.0011158264,0.00032213522,0.00025277282,0.00039361534,0.035379965,0.050165046,0.004707364,0.035739116,0.86981547],"study_design_scores_gemma":[0.00016308484,0.0009709016,0.0048815836,0.00017642944,0.00072936295,0.000672077,0.0005392967,0.7267625,0.12168091,0.022620382,0.12067318,0.00013036413],"about_ca_topic_score_codex":0.0019375908,"about_ca_topic_score_gemma":0.0043482883,"teacher_disagreement_score":0.006526666,"about_ca_system_score_codex":0.0005501505,"about_ca_system_score_gemma":0.0007433454,"threshold_uncertainty_score":0.021833897},"labels":[],"label_agreement":null},{"id":"W7111684846","doi":"","title":"MetaAligner::Towards Generalizable Multi-Objective Alignment of Language Models","year":2024,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Inference; Preference; Language model; Focus (optics); Training set; Machine translation","score_opus":0.12239382246347343,"score_gpt":0.31802296009807846,"score_spread":0.19562913763460504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7111684846","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005943498,0.00041125066,0.9805207,0.00019428402,0.00008242571,0.000094219424,0.00033112537,0.010990283,0.0014321433],"genre_scores_gemma":[0.18782091,0.0005584712,0.7970866,0.00084699114,0.00012387833,0.00051736366,0.0031823919,0.004064445,0.005798905],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975625,0.0009252743,0.00014131567,0.0008002764,0.0004204547,0.0001501123],"domain_scores_gemma":[0.9977998,0.0011955379,0.00017185997,0.00041874676,0.00031945726,0.00009454764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00302638,0.0026857671,0.0015471285,0.0011969397,0.00065611204,0.0018533145,0.002374739,0.0017286128,0.0064902618],"category_scores_gemma":[0.009049831,0.0013180928,0.0023720553,0.0011299552,0.0008940916,0.0037541864,0.002651179,0.004732026,0.0037461002],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026956655,0.00019418873,0.0017396401,0.00049967994,0.0003013341,0.00022236323,0.0003947025,0.44822857,0.008697036,0.018599879,0.013712042,0.50714105],"study_design_scores_gemma":[0.000026488342,0.000048105558,0.00016200518,0.00002325978,0.000022577584,0.00003640029,0.000039352388,0.97634345,0.0026587478,0.0171957,0.003423568,0.000020359796],"about_ca_topic_score_codex":0.0057765236,"about_ca_topic_score_gemma":0.00903454,"teacher_disagreement_score":0.0064902618,"about_ca_system_score_codex":0.0011794459,"about_ca_system_score_gemma":0.002587583,"threshold_uncertainty_score":0.021712065},"labels":[],"label_agreement":null},{"id":"W7112000808","doi":"","title":"Improve Chinese Clinical Named Entity Recognition Performance by Using the Graphical and Phonetic Feature","year":2018,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Feature (linguistics); Named-entity recognition; Pinyin; Entity linking; Named entity; Word (group theory)","score_opus":0.09316927895068974,"score_gpt":0.32582543036546285,"score_spread":0.2326561514147731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7112000808","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41639787,0.0037215177,0.5192689,0.0034652038,0.00092236523,0.0006083023,0.006956239,0.026393421,0.022266187],"genre_scores_gemma":[0.77230793,0.00074116944,0.21342358,0.00027762263,0.00020654353,0.00010381165,0.008677789,0.00036759913,0.003893992],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979634,0.0007600384,0.00022831015,0.00048092354,0.00040595984,0.00016145901],"domain_scores_gemma":[0.9945799,0.0028644414,0.00029341585,0.0008863333,0.001225051,0.00015087976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033899718,0.0012711443,0.0008069651,0.0025903978,0.00052290095,0.0018790183,0.00072004163,0.0007108811,0.0035247328],"category_scores_gemma":[0.01013365,0.00024521208,0.0012285864,0.0029100801,0.0002687715,0.0037160164,0.0012222676,0.0007909367,0.0033849012],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005056141,0.00037484607,0.02766761,0.0003380341,0.00029646629,0.00026523718,0.0004479172,0.010480969,0.034559045,0.0021720335,0.019367356,0.9035248],"study_design_scores_gemma":[0.00011177153,0.0005503031,0.05868281,0.00010083753,0.00083486474,0.0009773626,0.00081997894,0.7580548,0.1346905,0.005529126,0.039356776,0.0002908679],"about_ca_topic_score_codex":0.0069479924,"about_ca_topic_score_gemma":0.007421436,"teacher_disagreement_score":0.0069479924,"about_ca_system_score_codex":0.000464625,"about_ca_system_score_gemma":0.00092795945,"threshold_uncertainty_score":0.017928064},"labels":[],"label_agreement":null},{"id":"W7112701511","doi":"","title":"A latent concept topic model for robust topic inference using word embeddings","year":2016,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Topic model; Inference; Probabilistic latent semantic analysis; Word (group theory); Similarity (geometry); Latent Dirichlet allocation; Word embedding","score_opus":0.25338344721631956,"score_gpt":0.3337876138873781,"score_spread":0.08040416667105854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7112701511","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050313575,0.00032886694,0.9933789,0.00017157171,0.00004540543,0.0000499066,0.0002834656,0.0004341766,0.00027624078],"genre_scores_gemma":[0.2836808,0.0012956535,0.7029902,0.00031237816,0.00055432256,0.0011257713,0.005237278,0.00037622696,0.0044274074],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99782807,0.0009304509,0.00015520402,0.0006114905,0.00033117895,0.00014353219],"domain_scores_gemma":[0.9948085,0.0036307252,0.00034401185,0.00053833635,0.0005637954,0.00011464728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046280287,0.0011279731,0.0014176435,0.002462475,0.0007175366,0.0019621023,0.0023428882,0.001643817,0.0025925953],"category_scores_gemma":[0.013252932,0.00090926385,0.0019821906,0.0032536106,0.000839852,0.0042974,0.0020129404,0.0034946816,0.002108627],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053727557,0.0003580184,0.0055217966,0.0005440294,0.00056940166,0.0003277041,0.0010654966,0.3476247,0.007675678,0.109292105,0.015330911,0.511153],"study_design_scores_gemma":[0.000019978097,0.000026452928,0.00027200842,0.000021064223,0.000023701443,0.00004416297,0.000035847752,0.9685156,0.0004971679,0.028947894,0.001579923,0.000016178139],"about_ca_topic_score_codex":0.005737766,"about_ca_topic_score_gemma":0.0072955224,"teacher_disagreement_score":0.005737766,"about_ca_system_score_codex":0.00096172496,"about_ca_system_score_gemma":0.0017671065,"threshold_uncertainty_score":0.024475694},"labels":[],"label_agreement":null},{"id":"W7113040493","doi":"","title":"Generating Medical Instructions with Conditional Transformer","year":2023,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"Engineering and Physical Sciences Research Council; Nuffield Foundation; UK Research and Innovation","keywords":"Vocabulary; Language model; Transformer; Task (project management); Synthetic data; Training set","score_opus":0.09492307417161791,"score_gpt":0.2982224254437973,"score_spread":0.20329935127217938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7113040493","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10581444,0.00064865587,0.79712385,0.0015237258,0.00038820616,0.0004949701,0.017886808,0.06697559,0.0091437325],"genre_scores_gemma":[0.6455894,0.0003848422,0.29693532,0.0006568139,0.00013194761,0.0007062461,0.04173393,0.0019467677,0.011914769],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99964046,0.00009720442,0.000028931494,0.00013510603,0.00006634135,0.00003192383],"domain_scores_gemma":[0.9986155,0.0008441779,0.00007254771,0.00019781242,0.00022148382,0.000048400718],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00049624237,0.0008501294,0.00035337155,0.00063577417,0.00021485578,0.00054313,0.0010340525,0.00075415656,0.006016049],"category_scores_gemma":[0.0032407609,0.0002880682,0.0009606924,0.0004479764,0.00040808535,0.0010396211,0.00091839925,0.0010630591,0.0033505016],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012880208,0.00030400697,0.006477319,0.0008971704,0.00011885446,0.001011603,0.0006379194,0.32407486,0.03191486,0.02216817,0.089620896,0.52148634],"study_design_scores_gemma":[0.00005651475,0.000080104335,0.00054020714,0.000018526647,0.000024948999,0.00017638903,0.000038985672,0.9658698,0.013939516,0.008602153,0.010627992,0.000024908772],"about_ca_topic_score_codex":0.0051893587,"about_ca_topic_score_gemma":0.0077221105,"teacher_disagreement_score":0.006016049,"about_ca_system_score_codex":0.00072561577,"about_ca_system_score_gemma":0.0012946732,"threshold_uncertainty_score":0.020125687},"labels":[],"label_agreement":null},{"id":"W7116141524","doi":"10.1007/s10489-025-07044-6","title":"CMCTS: A Constrained Monte Carlo Tree Search framework for mathematical reasoning in large language model","year":2025,"lang":"en","type":"article","venue":"Applied Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Huawei Technologies (Canada)","funders":"","keywords":"Monte Carlo tree search; Tree (set theory); Action (physics); Action selection; Monte Carlo method; Baseline (sea); Selection (genetic algorithm)","score_opus":0.028245189370860593,"score_gpt":0.3237130342735292,"score_spread":0.29546784490266864,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7116141524","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011421989,0.000108815475,0.9963935,0.00015150542,0.000027293703,0.00005098873,0.00018498948,0.0010023066,0.000938512],"genre_scores_gemma":[0.13033223,0.00036380402,0.86249435,0.0003739065,0.00017611164,0.00059825956,0.0013243925,0.0011270343,0.003209944],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99769956,0.0012034879,0.000116631745,0.00025618018,0.0006010577,0.0001230524],"domain_scores_gemma":[0.9904246,0.0076736403,0.00026305587,0.00064426055,0.00067371014,0.00032071577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038379852,0.0010664759,0.0018924212,0.0020299098,0.0011001822,0.0029426743,0.0043990375,0.0025444091,0.013979471],"category_scores_gemma":[0.022909792,0.0011301371,0.0021486925,0.002660032,0.0016497838,0.0036713402,0.0032528176,0.003512361,0.0025424627],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022516774,0.00013318112,0.00070241955,0.00023311752,0.00018301947,0.00015657741,0.00021120328,0.6136288,0.0007482368,0.26839072,0.011502253,0.10388529],"study_design_scores_gemma":[0.000017287637,0.0000065658164,0.00001786562,0.000008821969,0.000008602704,0.000009181377,0.000007272043,0.9324458,0.00011500481,0.066134624,0.001223829,0.0000051162237],"about_ca_topic_score_codex":0.01580671,"about_ca_topic_score_gemma":0.02337698,"teacher_disagreement_score":0.01580671,"about_ca_system_score_codex":0.0018191438,"about_ca_system_score_gemma":0.0041153296,"threshold_uncertainty_score":0.046765983},"labels":[],"label_agreement":null},{"id":"W7117259628","doi":"10.1145/3786333","title":"A Survey on Large Language Models for Mathematical Reasoning","year":2025,"lang":"en","type":"article","venue":"ACM Computing Surveys","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"National Natural Science Foundation of China","keywords":"Cognition; Verbal reasoning; Automated reasoning; Psychology of reasoning; Reinforcement learning; Case-based reasoning; Language model; Qualitative reasoning","score_opus":0.04404501978785593,"score_gpt":0.32054745612532104,"score_spread":0.27650243633746513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117259628","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053901486,0.27074036,0.67641604,0.011175104,0.0009586585,0.00021320664,0.0023225206,0.0026975826,0.030086406],"genre_scores_gemma":[0.122405276,0.45075625,0.39138788,0.004373375,0.0052642226,0.0009336737,0.008741168,0.0015598693,0.014578252],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973477,0.0011053004,0.00026893607,0.00037546075,0.0007739238,0.00012867921],"domain_scores_gemma":[0.9915336,0.0066621685,0.00022841213,0.00085119554,0.0005921171,0.00013268663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038237406,0.00173179,0.0021288018,0.0033056038,0.00073601,0.004069734,0.0031875777,0.0018649625,0.01392332],"category_scores_gemma":[0.0138830105,0.0012942293,0.002367716,0.005619506,0.0013577673,0.0099572865,0.002452202,0.003990862,0.004986988],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010169677,0.0001955954,0.0015027365,0.0040010903,0.0002131361,0.00020728658,0.00047440655,0.035344888,0.0008583919,0.39952222,0.062352464,0.49522614],"study_design_scores_gemma":[0.00003105366,0.00007217356,0.0009476578,0.001297931,0.0001138517,0.0003919089,0.00017896704,0.13862383,0.00070862734,0.56441104,0.2931411,0.000081901526],"about_ca_topic_score_codex":0.0040180488,"about_ca_topic_score_gemma":0.0039813304,"teacher_disagreement_score":0.01392332,"about_ca_system_score_codex":0.0024564841,"about_ca_system_score_gemma":0.002678056,"threshold_uncertainty_score":0.04657811},"labels":[],"label_agreement":null},{"id":"W7117277568","doi":"","title":"Can LLMs Predict Their Own Failures? Self-Awareness via Internal Circuits","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Correctness; Inference; Process (computing); Decoding methods; Sequence (biology); Internal model","score_opus":0.02677271643426992,"score_gpt":0.2505918237108248,"score_spread":0.22381910727655485,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117277568","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4845942,0.00079519907,0.4879492,0.0015516124,0.0001434236,0.00009228897,0.0012673478,0.018510351,0.005096299],"genre_scores_gemma":[0.9646721,0.00007176549,0.032602105,0.00021553281,0.000021156458,0.00003920225,0.00088072504,0.0005248332,0.0009726882],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9987447,0.00035756014,0.000067872956,0.00040731387,0.00026314566,0.00015934462],"domain_scores_gemma":[0.9925558,0.0037931723,0.0005484851,0.002120652,0.0007090746,0.00027274684],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023826985,0.00086888665,0.00074414344,0.000797711,0.00046169234,0.0018115668,0.0017724776,0.0010152197,0.0025861682],"category_scores_gemma":[0.019445006,0.00054235343,0.0006467976,0.00044058613,0.0013593731,0.004778546,0.0018701766,0.0017565305,0.001143064],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013504605,0.0002385757,0.073186174,0.00041395862,0.00023696461,0.0006004085,0.0014116546,0.47506648,0.0332403,0.028851466,0.0152661195,0.3701374],"study_design_scores_gemma":[0.000018706274,0.000066251676,0.0020884643,0.000020309382,0.0000298337,0.000065836044,0.00008722091,0.96102184,0.009077489,0.026314566,0.0011887507,0.000020712516],"about_ca_topic_score_codex":0.0043484704,"about_ca_topic_score_gemma":0.0052720807,"teacher_disagreement_score":0.0043484704,"about_ca_system_score_codex":0.0010072442,"about_ca_system_score_gemma":0.0011276309,"threshold_uncertainty_score":0.012601078},"labels":[],"label_agreement":null},{"id":"W7117339582","doi":"","title":"Generalization of RLVR Using Causal Reasoning as a Testbed","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Division of Chemistry; Institute for Catastrophic Loss Reduction; National Science Foundation","keywords":"Generalization; Probabilistic logic; Counterfactual thinking; Causal model; Inference; Semantics (computer science); Causal structure; Construct (python library); Bayesian network; Causal inference","score_opus":0.056204327718486,"score_gpt":0.3069614930356174,"score_spread":0.25075716531713144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117339582","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3920688,0.001405838,0.5646923,0.0023967992,0.000248807,0.0004897377,0.0021198813,0.024800148,0.011777629],"genre_scores_gemma":[0.87521064,0.00014961617,0.11904929,0.00065342616,0.00003478652,0.0002864381,0.0024812366,0.00074053224,0.0013941386],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99451905,0.002873166,0.00026631108,0.0013759672,0.00059433567,0.0003712002],"domain_scores_gemma":[0.9529312,0.035356194,0.0011543719,0.008526409,0.001454296,0.0005774728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010372801,0.0017539598,0.0011099682,0.0009490073,0.00061825564,0.001536327,0.002579309,0.002065883,0.003997025],"category_scores_gemma":[0.0625914,0.00091327395,0.0017321793,0.00061162154,0.0017592979,0.0040666135,0.002901982,0.0056589344,0.00092990935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005478135,0.00052325544,0.0076965415,0.000400391,0.0002078524,0.00019220747,0.00037350898,0.88136584,0.003606168,0.011474289,0.006820488,0.08679161],"study_design_scores_gemma":[0.000051695974,0.000085940155,0.00041062714,0.000019284593,0.000017000668,0.000023554563,0.000031476382,0.9864136,0.0014376197,0.010828251,0.0006687934,0.000012172863],"about_ca_topic_score_codex":0.013907257,"about_ca_topic_score_gemma":0.015320354,"teacher_disagreement_score":0.013907257,"about_ca_system_score_codex":0.002626039,"about_ca_system_score_gemma":0.002004327,"threshold_uncertainty_score":0.054857254},"labels":[],"label_agreement":null},{"id":"W7117481005","doi":"10.2139/ssrn.5978094","title":"A Comprehensive Survey on Large Language Models: From Pre-training to Autonomous Agents","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Canada West","funders":"","keywords":"Interoperability; Transformative learning; Architecture; Key (lock); Natural language understanding; Inference; Open research","score_opus":0.05460445667758154,"score_gpt":0.3138997307637531,"score_spread":0.2592952740861716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117481005","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010793804,0.07755289,0.8935498,0.005352769,0.00060364575,0.00015819441,0.0015144245,0.0044038966,0.006070536],"genre_scores_gemma":[0.27106598,0.11819331,0.575168,0.003436334,0.003802643,0.0010227313,0.011553112,0.0025734878,0.013184427],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99708754,0.0014710336,0.00022458912,0.00062614965,0.00045872052,0.00013211579],"domain_scores_gemma":[0.98054767,0.016234359,0.00028201347,0.0016913654,0.0009204994,0.0003240159],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004449454,0.0017006965,0.0027828973,0.0019111963,0.0006973682,0.0030769652,0.0031484684,0.0021681872,0.004713512],"category_scores_gemma":[0.02244522,0.0012428287,0.0015337301,0.0027532855,0.0009159964,0.006792081,0.0025048603,0.0034574876,0.0037468104],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022571141,0.00035407298,0.0023632054,0.0020322048,0.0002907474,0.00009492194,0.0003229444,0.09038093,0.001663199,0.028704561,0.035636492,0.83793104],"study_design_scores_gemma":[0.00005516462,0.00025258577,0.0013064403,0.0006017382,0.00019607531,0.00022023739,0.0002333507,0.79746306,0.0027930941,0.14092603,0.05588569,0.00006658282],"about_ca_topic_score_codex":0.005977315,"about_ca_topic_score_gemma":0.007265405,"teacher_disagreement_score":0.005977315,"about_ca_system_score_codex":0.0013876349,"about_ca_system_score_gemma":0.0034039677,"threshold_uncertainty_score":0.023531258},"labels":[],"label_agreement":null},{"id":"W7117484981","doi":"10.1016/j.eswa.2025.130958","title":"Topic modeling and alignment with large language models for multi-labeled text corpora","year":2025,"lang":"en","type":"article","venue":"Expert Systems with Applications","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Fundamental Research Funds for the Central Universities; China Postdoctoral Science Foundation; National Natural Science Foundation of China","keywords":"Interpretability; Topic model; Language model; Probabilistic logic; Coherence (philosophical gambling strategy); Semantics (computer science); Latent Dirichlet allocation","score_opus":0.03005171954682459,"score_gpt":0.28740508879336835,"score_spread":0.25735336924654373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117484981","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0069574653,0.0009961647,0.98617315,0.00039189428,0.00018224203,0.00013386089,0.0009011235,0.003718505,0.00054555445],"genre_scores_gemma":[0.17280562,0.0014668189,0.80311036,0.00032028966,0.00096263684,0.001142686,0.014348002,0.0018221876,0.004021417],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99107224,0.0047263103,0.00067101774,0.0023213825,0.00084384566,0.0003652488],"domain_scores_gemma":[0.9786129,0.01610253,0.00083223975,0.0024049364,0.0016842886,0.00036309232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008410086,0.0019815264,0.0027723417,0.004875924,0.0022971751,0.0040465146,0.0033520544,0.0028450976,0.0038397338],"category_scores_gemma":[0.028687429,0.0018559756,0.0033729062,0.006107728,0.0009185768,0.006576587,0.0029114238,0.005273844,0.005154175],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011528942,0.0006599505,0.004152684,0.0009898784,0.0009967373,0.0005628041,0.001311676,0.1466841,0.013194799,0.029173302,0.03202434,0.76909673],"study_design_scores_gemma":[0.00006713939,0.000077854915,0.0010732252,0.000049876013,0.00016782696,0.00016846131,0.00019338643,0.94695824,0.0041518165,0.040558994,0.0064706793,0.00006239522],"about_ca_topic_score_codex":0.0069727805,"about_ca_topic_score_gemma":0.012629913,"teacher_disagreement_score":0.008410086,"about_ca_system_score_codex":0.0013860473,"about_ca_system_score_gemma":0.0028858904,"threshold_uncertainty_score":0.044477284},"labels":[],"label_agreement":null},{"id":"W7117513518","doi":"10.1109/icprs66293.2025.11302825","title":"Embedding Confidence to Enhance Trust in AI Document Entity Extraction","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Springboard (Canada)","funders":"","keywords":"Flagging; Reliability (semiconductor); Workflow; Embedding; Matching (statistics); Information extraction; Data extraction","score_opus":0.013574009715650467,"score_gpt":0.35834340442342955,"score_spread":0.34476939470777906,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117513518","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.120023325,0.002151256,0.8605075,0.0018496132,0.00015097261,0.00025151798,0.0021494906,0.009279287,0.0036370088],"genre_scores_gemma":[0.8051624,0.00040091597,0.18963523,0.00019625176,0.00012522055,0.00011795624,0.0024977908,0.00052712107,0.0013372657],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9871197,0.005426324,0.0016121,0.0022415929,0.0030745764,0.00052575703],"domain_scores_gemma":[0.88256127,0.07743977,0.011331574,0.017617587,0.01020679,0.00084303727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013223357,0.001056918,0.0011601312,0.00454696,0.0009788194,0.0047776704,0.0017418069,0.0018585772,0.0024134342],"category_scores_gemma":[0.15486544,0.0006552088,0.00093444384,0.0030198453,0.0011757694,0.010297263,0.00536916,0.0024522077,0.0017268141],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011161,0.00023948844,0.052222148,0.001309528,0.00032095908,0.000516523,0.003826812,0.07681642,0.016655285,0.025086729,0.012583115,0.80930686],"study_design_scores_gemma":[0.000055880104,0.00023017867,0.008060283,0.00021980755,0.00014358255,0.0004738136,0.00065089855,0.89807993,0.027958192,0.051713414,0.0122920815,0.00012187831],"about_ca_topic_score_codex":0.0031534801,"about_ca_topic_score_gemma":0.004099855,"teacher_disagreement_score":0.013223357,"about_ca_system_score_codex":0.001340381,"about_ca_system_score_gemma":0.0019912024,"threshold_uncertainty_score":0.06993258},"labels":[],"label_agreement":null},{"id":"W7117578138","doi":"10.1016/j.jss.2025.112763","title":"LeCov: Multi-level testing criteria for large language models","year":2025,"lang":"en","type":"article","venue":"Journal of Systems and Software","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Mila - Quebec Artificial Intelligence Institute","funders":"Japan Science and Technology Agency; Japan Society for the Promotion of Science; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research","keywords":"Set (abstract data type); Trustworthiness; Prioritization; Test (biology); Test strategy; Risk-based testing; Non-regression testing","score_opus":0.08902155403723257,"score_gpt":0.32290924107423324,"score_spread":0.2338876870370007,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117578138","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017734284,0.00019889603,0.9752514,0.00041483983,0.00005406526,0.000121851066,0.0004505707,0.0034715866,0.0023024445],"genre_scores_gemma":[0.6342886,0.00014134713,0.35482442,0.0003666867,0.00014592259,0.00053209614,0.0024657575,0.0023558396,0.0048793745],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9912411,0.0039384984,0.00060497614,0.000706153,0.0026006626,0.00090857816],"domain_scores_gemma":[0.9284092,0.05580137,0.001826196,0.0041001844,0.007887809,0.0019752253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009532462,0.0016516823,0.0016725584,0.004069523,0.0011166316,0.0028793132,0.003956275,0.0023262957,0.011256952],"category_scores_gemma":[0.07309415,0.0008339647,0.002012988,0.0018578997,0.0019148097,0.0060487026,0.00462452,0.0033760332,0.0015063781],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019319555,0.00043757338,0.009438284,0.0009860976,0.00031855688,0.0008069133,0.0003958233,0.38475943,0.01135063,0.2503955,0.024911903,0.31426734],"study_design_scores_gemma":[0.00005483342,0.00010638468,0.00037671297,0.000040972685,0.000026651598,0.00006384155,0.000031484342,0.8993604,0.0029548595,0.09590202,0.0010592233,0.000022580682],"about_ca_topic_score_codex":0.0039235204,"about_ca_topic_score_gemma":0.004439214,"teacher_disagreement_score":0.011256952,"about_ca_system_score_codex":0.0021953168,"about_ca_system_score_gemma":0.0027977233,"threshold_uncertainty_score":0.05041313},"labels":[],"label_agreement":null},{"id":"W7117631606","doi":"","title":"A Comedy of Estimators: On KL Regularization in RL Training of LLMs","year":2025,"lang":"","type":"article","venue":"ArXiv.org","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Lawrence Livermore National Laboratory; Fonds de recherche du Québec – Nature et technologies; Office of Science; Natural Sciences and Engineering Research Council of Canada; Canadian Institute for Advanced Research; Laboratory Directed Research and Development; U.S. Department of Energy","keywords":"Regularization (linguistics); Estimator; Divergence (linguistics); Reinforcement learning; Asynchronous communication","score_opus":0.06209032547014128,"score_gpt":0.2915530558813591,"score_spread":0.2294627304112178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117631606","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020888554,0.0007906464,0.973912,0.00083979394,0.000077138335,0.00007172846,0.00004677798,0.0013473125,0.0020260508],"genre_scores_gemma":[0.64315003,0.0006490426,0.350246,0.0014514804,0.00013694228,0.00041513453,0.00026122463,0.0010617228,0.0026283558],"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99458486,0.0032464045,0.00027188266,0.00072367827,0.0008430046,0.00033022813],"domain_scores_gemma":[0.9727729,0.022185257,0.00094431493,0.002190716,0.0014999375,0.00040685284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00943846,0.0017509392,0.0017076463,0.0008085342,0.0007855016,0.002270418,0.0024561465,0.0025963464,0.0027235562],"category_scores_gemma":[0.06802106,0.0011228166,0.00080384803,0.0006182984,0.003224787,0.0038797818,0.0036220762,0.0055336663,0.0010333896],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000417045,0.00018174919,0.0036091392,0.00028528608,0.00014192217,0.00015992738,0.0003381273,0.8220253,0.0037270773,0.04768782,0.004158793,0.11726777],"study_design_scores_gemma":[0.00003607781,0.00006362858,0.00016499417,0.000055661138,0.000012526425,0.00002486378,0.000026054857,0.98014694,0.0017702853,0.017034745,0.0006485833,0.00001556136],"about_ca_topic_score_codex":0.005417113,"about_ca_topic_score_gemma":0.0065444065,"teacher_disagreement_score":0.00943846,"about_ca_system_score_codex":0.0017208124,"about_ca_system_score_gemma":0.0026622366,"threshold_uncertainty_score":0.04991591},"labels":[],"label_agreement":null},{"id":"W7118165281","doi":"10.1109/aiccsa66935.2025.11315432","title":"A Novel Competency Tagging Method Through Semantic Search Using Fine-Tuned LLM","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Université TÉLUQ","funders":"","keywords":"Matching (statistics); Semantic matching; Embedding; Similarity (geometry); Adaptation (eye); Semantic similarity; Training set; Process (computing)","score_opus":0.09893633220898373,"score_gpt":0.3777495625801211,"score_spread":0.27881323037113737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7118165281","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039956257,0.0011000162,0.9439564,0.00040907777,0.00016815549,0.00023812623,0.0011389364,0.01014537,0.002887791],"genre_scores_gemma":[0.36911133,0.00054926163,0.613419,0.0007934858,0.00021534928,0.00039275378,0.0071502184,0.0009837538,0.007384912],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984871,0.00044446552,0.00015694223,0.0003864232,0.00037583982,0.00014911075],"domain_scores_gemma":[0.9985856,0.0004967053,0.0001375035,0.00029710925,0.00038986362,0.00009328375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013066399,0.0015037468,0.0013302275,0.004136945,0.0008246149,0.0013901197,0.0015198266,0.00139273,0.0034637314],"category_scores_gemma":[0.005384458,0.0004060318,0.0015355671,0.0024974593,0.00068413746,0.0034122923,0.002046987,0.0013417521,0.0045686155],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045562763,0.0006557776,0.0059481296,0.0006265463,0.00019878757,0.0004828035,0.00048912165,0.045834217,0.055743065,0.008984269,0.02756393,0.8530178],"study_design_scores_gemma":[0.00011230413,0.00024850998,0.0020390414,0.000058328773,0.000112757865,0.00055726676,0.00038485814,0.941294,0.027918888,0.015143163,0.012036686,0.000094324016],"about_ca_topic_score_codex":0.005994636,"about_ca_topic_score_gemma":0.0109168375,"teacher_disagreement_score":0.005994636,"about_ca_system_score_codex":0.00087264145,"about_ca_system_score_gemma":0.0023348716,"threshold_uncertainty_score":0.011919498},"labels":[],"label_agreement":null},{"id":"W7118177928","doi":"10.1109/aiccsa66935.2025.11315347","title":"Evaluating Arabic Language Embedding Models for Semantic Retrieval in Fatwas","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Generative grammar; Sentence; Embedding; Natural language; Universal Networking Language; Context (archaeology); Transformer; Language model","score_opus":0.07815616242647719,"score_gpt":0.39354135309769683,"score_spread":0.3153851906712196,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7118177928","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9044237,0.0006800211,0.082766816,0.00051271246,0.00008337138,0.0003197563,0.0007805423,0.0035425622,0.0068904567],"genre_scores_gemma":[0.9602457,0.0002081712,0.03563233,0.00005787557,0.000024098923,0.00011111224,0.0019035287,0.00012304235,0.0016940398],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987192,0.0008693691,0.00006943531,0.000173938,0.00012143281,0.000046597277],"domain_scores_gemma":[0.9921774,0.0065130284,0.00018034279,0.0004730538,0.0005418341,0.00011436685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032724412,0.0008140834,0.00044201853,0.0011005328,0.0003347525,0.0015251806,0.00075624796,0.0008098303,0.0026018003],"category_scores_gemma":[0.010487097,0.00020671063,0.0005757073,0.0006192632,0.00042292353,0.0023549215,0.0009301264,0.00093754823,0.0012186639],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023301789,0.0012510633,0.015825182,0.00091723213,0.00044624077,0.00037103874,0.0028114077,0.47048998,0.016145846,0.0114780385,0.005485959,0.47244778],"study_design_scores_gemma":[0.000048311103,0.00025425057,0.001523014,0.000020452542,0.00007301956,0.000058653284,0.00041333892,0.98754126,0.006318105,0.002440086,0.0012881174,0.000021516042],"about_ca_topic_score_codex":0.0045558973,"about_ca_topic_score_gemma":0.00419851,"teacher_disagreement_score":0.0045558973,"about_ca_system_score_codex":0.00083051383,"about_ca_system_score_gemma":0.0007955374,"threshold_uncertainty_score":0.017306566},"labels":[],"label_agreement":null},{"id":"W7118191471","doi":"10.1109/aiccsa66935.2025.11315450","title":"Ensemble Hallucination Detection in Arabic Large Language Models: A Multi-Modal Linguistic and Self-Consistency Approach","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Abitibi-Témiscamingue","funders":"","keywords":"Readability; Arabic; Heuristics; Feature (linguistics); Software deployment","score_opus":0.020533798646119733,"score_gpt":0.26024224383999667,"score_spread":0.23970844519387693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7118191471","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.310012,0.00093893503,0.6738809,0.00064097566,0.00012992167,0.00017299058,0.0006798144,0.01155111,0.001993414],"genre_scores_gemma":[0.8744391,0.00016837275,0.12166642,0.00018531903,0.00006817519,0.00007734343,0.0016626512,0.0003578941,0.0013747119],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9990151,0.0003752626,0.000070757975,0.0002577185,0.00018186985,0.00009926595],"domain_scores_gemma":[0.9965364,0.0016437218,0.00024598755,0.0006001783,0.00076754316,0.00020618601],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002591529,0.0012373385,0.0008027198,0.0010539101,0.0005395933,0.001199441,0.0011721412,0.0006953792,0.0010126729],"category_scores_gemma":[0.008522578,0.0002909609,0.0008128192,0.00049301365,0.00038616278,0.0018216599,0.0020597843,0.0014746814,0.00094866334],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010286134,0.00043173193,0.024319222,0.00021893551,0.00045545044,0.00050844793,0.0010769573,0.12462558,0.044259083,0.0023871728,0.009094027,0.79159474],"study_design_scores_gemma":[0.000014899769,0.000091501715,0.0017216036,0.000012483677,0.000051131552,0.00012082443,0.00021120017,0.9844726,0.009942225,0.0023641086,0.00097576645,0.000021692125],"about_ca_topic_score_codex":0.0038076602,"about_ca_topic_score_gemma":0.005459826,"teacher_disagreement_score":0.0038076602,"about_ca_system_score_codex":0.00041847862,"about_ca_system_score_gemma":0.00072445837,"threshold_uncertainty_score":0.013705492},"labels":[],"label_agreement":null},{"id":"W7123334923","doi":"10.1109/slaai-icai68534.2025.11318480","title":"Task-Specific Knowledge Distillation for Accurate and Efficient Text Summarization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Automatic summarization; Inference; Distillation; Semantic similarity; Process (computing); Structured prediction; Natural language; Frame (networking)","score_opus":0.02829601060985867,"score_gpt":0.2848769062723144,"score_spread":0.25658089566245573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123334923","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03456995,0.0006144775,0.9474467,0.0007169487,0.00012787046,0.00015871343,0.0008673033,0.012591258,0.0029068151],"genre_scores_gemma":[0.50332683,0.00038171472,0.48398617,0.0005223251,0.00015295674,0.00031738894,0.0053310227,0.00061609125,0.005365512],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991979,0.00027023684,0.00006749726,0.00024131713,0.00014961575,0.000073415926],"domain_scores_gemma":[0.9984768,0.0007720513,0.0001013135,0.000385679,0.00018751241,0.00007665602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012293058,0.0010817734,0.00090505707,0.00085278443,0.00065468607,0.0011400855,0.0019511731,0.0012373286,0.0040401225],"category_scores_gemma":[0.005754157,0.00035611526,0.0010939462,0.00094242295,0.0005985042,0.0031397704,0.0021441244,0.002472589,0.0018034719],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004135877,0.00040478815,0.0010464637,0.00049115205,0.00013126867,0.00035680036,0.0006422468,0.2476737,0.03132292,0.022230364,0.016394177,0.67889255],"study_design_scores_gemma":[0.000030732517,0.000074080664,0.00017548016,0.0000115451985,0.0000265583,0.000049332408,0.000062010025,0.970072,0.009124743,0.015875416,0.0044821454,0.000016134652],"about_ca_topic_score_codex":0.003788127,"about_ca_topic_score_gemma":0.0066733365,"teacher_disagreement_score":0.0040401225,"about_ca_system_score_codex":0.0006988062,"about_ca_system_score_gemma":0.001348648,"threshold_uncertainty_score":0.013515532},"labels":[],"label_agreement":null},{"id":"W7123345670","doi":"10.1109/esem64174.2025.00054","title":"Interrogative Comments Posed by Review Comment Generators: An Empirical Study of Gerrit","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; University of Waterloo","funders":"","keywords":"Interrogative; Rhetorical question; Set (abstract data type); Interrogative word; Code (set theory); Empirical research; Similarity (geometry)","score_opus":0.06661327640412672,"score_gpt":0.3897362598836552,"score_spread":0.32312298347952845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7123345670","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9883044,0.00070378766,0.0042031356,0.001890925,0.00007187214,0.00045933225,0.00026773172,0.00046091346,0.003637887],"genre_scores_gemma":[0.9884306,0.0004354593,0.006133194,0.001717832,0.00006630289,0.0004934558,0.00042061586,0.00033070106,0.0019717303],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9088561,0.066594385,0.0037556582,0.005937638,0.013183457,0.0016728563],"domain_scores_gemma":[0.33444694,0.5184359,0.074517176,0.025711408,0.039215866,0.007672766],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08289317,0.00072197,0.00082680926,0.003948631,0.0025320554,0.0041126194,0.0025837289,0.002574724,0.0024308688],"category_scores_gemma":[0.39353028,0.0007691276,0.00053438736,0.002045505,0.0034630469,0.0052592508,0.0040977746,0.0025495393,0.0018556231],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011387285,0.0013006633,0.30331805,0.002417014,0.00018849027,0.0034843893,0.4984717,0.0010394427,0.008734927,0.0024770033,0.023229372,0.15420023],"study_design_scores_gemma":[0.00040591674,0.0033242234,0.44462335,0.0027576974,0.00029214428,0.0067671384,0.37318397,0.022908395,0.012439317,0.006783065,0.12571025,0.0008045442],"about_ca_topic_score_codex":0.0030453452,"about_ca_topic_score_gemma":0.005300223,"teacher_disagreement_score":0.9171068,"about_ca_system_score_codex":0.0027613398,"about_ca_system_score_gemma":0.003712662,"threshold_uncertainty_score":0.4383862},"labels":[],"label_agreement":null},{"id":"W7124295090","doi":"10.71781/34086","title":"Evaluating and improving mathematical reasoning in large language models via skill combinations","year":2025,"lang":"en","type":"dissertation","venue":"Open MIND","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Defense Advanced Research Projects Agency; Alliance de recherche numérique du Canada; National Science Foundation","keywords":"Context (archaeology); Subject (documents); Intentionality; Expected utility hypothesis","score_opus":0.04240423615353835,"score_gpt":0.3769850964572031,"score_spread":0.33458086030366474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7124295090","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2152321,0.0021207205,0.74038285,0.0015810919,0.00021702383,0.00065295724,0.0023111159,0.025832694,0.011669453],"genre_scores_gemma":[0.5859652,0.00058727025,0.4002482,0.00048533877,0.00007614744,0.0004887956,0.006091389,0.0008481783,0.0052094776],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950937,0.0023106223,0.00031481046,0.0011505057,0.0009300394,0.00020038527],"domain_scores_gemma":[0.98422885,0.012067023,0.0005409968,0.0015413415,0.0011689905,0.00045286672],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005282732,0.0021095183,0.001344115,0.0018787937,0.00071503787,0.004127944,0.0028518983,0.002003256,0.0069652563],"category_scores_gemma":[0.026598299,0.0009813701,0.002593434,0.0012668528,0.0008291621,0.0070508546,0.0035795423,0.002833595,0.0028960896],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015622638,0.0011566498,0.014164548,0.0015228261,0.0008133319,0.00047598963,0.0014023107,0.26995057,0.017549826,0.00897888,0.010200826,0.672222],"study_design_scores_gemma":[0.000097976095,0.00032817555,0.001887128,0.00008048018,0.0002236019,0.000089056855,0.0002870944,0.96715784,0.0077793654,0.016517164,0.0055091926,0.00004282945],"about_ca_topic_score_codex":0.008491139,"about_ca_topic_score_gemma":0.016027005,"teacher_disagreement_score":0.008491139,"about_ca_system_score_codex":0.0019181516,"about_ca_system_score_gemma":0.0021315687,"threshold_uncertainty_score":0.027938068},"labels":[],"label_agreement":null},{"id":"W7124297238","doi":"10.65109/wjjv5555","title":"SCMRAG: Self-Corrective Multihop Retrieval Augmented Generation System for LLM Agents","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Graph; Benchmark (surveying); Mechanism (biology); Similarity (geometry); Knowledge graph; Range (aeronautics); Matching (statistics); Interaction information","score_opus":0.0368575481277664,"score_gpt":0.2890033105762336,"score_spread":0.2521457624484672,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7124297238","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033343952,0.0007538936,0.77040434,0.00045417494,0.0001946894,0.00078480004,0.0018022141,0.18473835,0.007523601],"genre_scores_gemma":[0.21568854,0.00031230293,0.7659455,0.0005297921,0.000049305658,0.0005205188,0.006198776,0.001905966,0.008849342],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99944025,0.00013870215,0.000048332528,0.00014266407,0.00018869985,0.00004126608],"domain_scores_gemma":[0.99875057,0.00035480622,0.00009760666,0.00050787913,0.00020633261,0.00008274264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012876379,0.00085830136,0.00063352135,0.0012741976,0.0005875491,0.0011291858,0.0023194454,0.001183308,0.0057875267],"category_scores_gemma":[0.0037648699,0.00037704923,0.000719819,0.00073562335,0.00043962893,0.0019438659,0.0024383618,0.0010552362,0.0032707495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070309197,0.000544981,0.0032478618,0.00073552824,0.00019380722,0.00072173745,0.00094332616,0.03493311,0.03488987,0.0074676103,0.08935694,0.8262622],"study_design_scores_gemma":[0.00039731208,0.0004067406,0.0020318576,0.00007492475,0.00016004736,0.00059187633,0.00036322008,0.8271473,0.03799079,0.015426275,0.115265355,0.000144291],"about_ca_topic_score_codex":0.005415523,"about_ca_topic_score_gemma":0.0082573015,"teacher_disagreement_score":0.0057875267,"about_ca_system_score_codex":0.00058985216,"about_ca_system_score_gemma":0.000956761,"threshold_uncertainty_score":0.019361198},"labels":[],"label_agreement":null},{"id":"W7124425818","doi":"10.1093/jamiaopen/ooaf179","title":"MedSlice: fine-tuned large language models for secure clinical note sectioning","year":2025,"lang":"en","type":"article","venue":"JAMIA Open","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; McGill University Health Centre","funders":"","keywords":"Language model; Key (lock); Modeling language; Natural language; Action (physics)","score_opus":0.04370099868506119,"score_gpt":0.39791305372763885,"score_spread":0.3542120550425777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7124425818","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060637552,0.002854551,0.67033017,0.002099008,0.0006305999,0.0010874146,0.0263789,0.23228304,0.00369883],"genre_scores_gemma":[0.33906892,0.0011084172,0.59825516,0.0013451013,0.00023577538,0.0014308501,0.048081487,0.005386527,0.0050877314],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99817264,0.0006908724,0.00022169703,0.00056349783,0.00024958965,0.00010159733],"domain_scores_gemma":[0.99241805,0.005412971,0.0003972205,0.00089582265,0.000646379,0.00022946819],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042383657,0.0015399576,0.00064666814,0.0017557472,0.00051493774,0.001815351,0.0022383183,0.0014652978,0.0066572884],"category_scores_gemma":[0.020318596,0.0007445854,0.0020533262,0.000830297,0.00047372628,0.0021714629,0.0024198028,0.0024240077,0.0051255887],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025293934,0.00063771283,0.023035172,0.0014621473,0.00092572987,0.0010717814,0.0011589024,0.17364271,0.014440679,0.00669713,0.14098977,0.6334089],"study_design_scores_gemma":[0.00036215852,0.00022447403,0.0033596915,0.00020156926,0.00017591247,0.0004404017,0.00023393595,0.9330569,0.013096196,0.0144232055,0.034310136,0.00011546242],"about_ca_topic_score_codex":0.014379416,"about_ca_topic_score_gemma":0.02332769,"teacher_disagreement_score":0.014379416,"about_ca_system_score_codex":0.001609781,"about_ca_system_score_gemma":0.0032215454,"threshold_uncertainty_score":0.028591454},"labels":[],"label_agreement":null},{"id":"W7125515560","doi":"10.1109/icft66708.2025.11336532","title":"STAR-Net: Advancing Conversational AI through Semantic Correlation and Contextual Response Optimization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Correlation; Semantics (computer science); Feature (linguistics); Context (archaeology)","score_opus":0.01345631533480501,"score_gpt":0.2639869250660548,"score_spread":0.2505306097312498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125515560","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046562916,0.0002810931,0.98470205,0.00043083573,0.00015562374,0.00007395057,0.0002665184,0.0063006766,0.003133037],"genre_scores_gemma":[0.22316125,0.00044328035,0.7636214,0.00069626945,0.0003386021,0.00032560003,0.0015495187,0.0020470547,0.0078169685],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99661154,0.0018357509,0.0001276941,0.0006524148,0.00059135037,0.00018124466],"domain_scores_gemma":[0.9946656,0.003964891,0.00012995988,0.0005647912,0.00045779574,0.00021693524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003343897,0.0012430382,0.0014583715,0.0010578874,0.001154047,0.0023225844,0.0028296679,0.0014837282,0.009661689],"category_scores_gemma":[0.010665062,0.00074182963,0.0014604707,0.0011050317,0.001186953,0.004219517,0.003925323,0.0038557097,0.003731513],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014402437,0.0007960261,0.0016648952,0.00090275327,0.00042828158,0.0002358716,0.0014087089,0.19529785,0.016514396,0.14839153,0.046467472,0.5864519],"study_design_scores_gemma":[0.0000382707,0.00004386247,0.000067426714,0.00001758307,0.000038972412,0.000023228695,0.0000767161,0.9275606,0.0026415864,0.06294919,0.006526197,0.000016320348],"about_ca_topic_score_codex":0.00670797,"about_ca_topic_score_gemma":0.011355846,"teacher_disagreement_score":0.009661689,"about_ca_system_score_codex":0.0009861201,"about_ca_system_score_gemma":0.0022062673,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W7125587093","doi":"10.1109/cascon66301.2025.00082","title":"From Requirements to Models: A Case Study of Large Language Models on Epidemiological Modeling","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Wellcome Trust","keywords":"Representation (politics); Modeling language; Key (lock); Language model; Data modeling; Transmission (telecommunications)","score_opus":0.18615072940077862,"score_gpt":0.395080489789336,"score_spread":0.2089297603885574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125587093","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2937734,0.0003807365,0.67471594,0.0049251113,0.000063875814,0.00072624645,0.000999508,0.0024318348,0.021983331],"genre_scores_gemma":[0.63081837,0.00034326332,0.3620967,0.0004798351,0.000030077337,0.000448836,0.00087355723,0.00091327995,0.0039960924],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98948425,0.007860554,0.00044901075,0.00048161924,0.0014044057,0.00032017208],"domain_scores_gemma":[0.9231902,0.06818462,0.0013890082,0.004055251,0.0025821573,0.0005987942],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0107582845,0.0008286487,0.00047481037,0.0010038222,0.0013261227,0.0027260273,0.00198071,0.0021020093,0.004196212],"category_scores_gemma":[0.0397827,0.0006567436,0.001398228,0.0011140498,0.0018415782,0.004137865,0.0033837107,0.0024789765,0.00065466575],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066377927,0.00087063457,0.012982869,0.0014475732,0.0001390211,0.010059679,0.021239366,0.47709098,0.01816427,0.35834703,0.009290769,0.089703985],"study_design_scores_gemma":[0.00014416616,0.0003289062,0.0020460803,0.0003292191,0.00010169333,0.0018840702,0.0053687273,0.79675525,0.018148499,0.09301097,0.081780404,0.00010195999],"about_ca_topic_score_codex":0.007810336,"about_ca_topic_score_gemma":0.008203619,"teacher_disagreement_score":0.0107582845,"about_ca_system_score_codex":0.0026176388,"about_ca_system_score_gemma":0.0026183869,"threshold_uncertainty_score":0.05689591},"labels":[],"label_agreement":null},{"id":"W7125589449","doi":"10.1109/cascon66301.2025.00096","title":"Streamlining Epidemiological Model Generation Using Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Representation (politics); Modeling language; Key (lock); Language model; Data modeling; Transmission (telecommunications)","score_opus":0.15714932206460852,"score_gpt":0.36338280448961313,"score_spread":0.2062334824250046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125589449","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014102219,0.000051000574,0.9766237,0.00031498025,0.000026028529,0.00021638905,0.00042901805,0.0070926542,0.0011440109],"genre_scores_gemma":[0.1748437,0.00017483176,0.8195741,0.00019270286,0.000021079817,0.0005176262,0.0016754655,0.0016484269,0.0013520456],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.99721587,0.0015013851,0.00022128627,0.0003176358,0.00063053204,0.00011337433],"domain_scores_gemma":[0.9765246,0.018688012,0.00074598426,0.002335341,0.0014956591,0.00021039988],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048998855,0.0012508134,0.0007629345,0.0011983869,0.0005948421,0.0022665954,0.0022067155,0.0010364254,0.005591929],"category_scores_gemma":[0.029270716,0.0010880848,0.0016808833,0.0006780787,0.00072480575,0.0029089125,0.002943091,0.002036578,0.0016806154],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045521674,0.00034395198,0.006181358,0.0012576444,0.00019015785,0.001165337,0.002198728,0.65583396,0.03191188,0.074016914,0.009348778,0.21709615],"study_design_scores_gemma":[0.00004370597,0.000054703683,0.00014431964,0.000052804906,0.000033703727,0.000088471046,0.0001583204,0.9586899,0.009992635,0.021610266,0.009109749,0.000021389342],"about_ca_topic_score_codex":0.0031520023,"about_ca_topic_score_gemma":0.0062635005,"teacher_disagreement_score":0.005591929,"about_ca_system_score_codex":0.00090550823,"about_ca_system_score_gemma":0.0023657042,"threshold_uncertainty_score":0.025913417},"labels":[],"label_agreement":null},{"id":"W7125593338","doi":"10.1109/cascon66301.2025.00027","title":"Designing Large Language Models for Specific Domains: A Case Study on Live Microbe Foods for Precision Nutrition","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Michael's Hospital; University of Toronto; Institute for Clinical Evaluative Sciences","funders":"Canadian Institutes of Health Research","keywords":"Variety (cybernetics); Generalization; Generative grammar; Function (biology); Process (computing); Domain (mathematical analysis)","score_opus":0.05672739262348838,"score_gpt":0.32485564401958866,"score_spread":0.2681282513961003,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125593338","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5522113,0.00097971,0.4236567,0.0043548984,0.00009549837,0.00083434896,0.0018920955,0.0048553054,0.011120117],"genre_scores_gemma":[0.5843926,0.0004324612,0.40431803,0.0005535297,0.00004483129,0.00050523056,0.0019073933,0.00081799336,0.007027908],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99621177,0.0029749214,0.00013174821,0.0003919361,0.00021034821,0.00007933914],"domain_scores_gemma":[0.96460265,0.032779865,0.00037965318,0.001228308,0.0006945323,0.00031492356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051421206,0.0009336829,0.000507341,0.00058901566,0.0010040594,0.0016876679,0.0019688527,0.002450355,0.0038927123],"category_scores_gemma":[0.024557365,0.0005092286,0.0008922032,0.0007420643,0.0010553212,0.004388066,0.0015868284,0.0018997742,0.0013049062],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017700226,0.003836551,0.045805365,0.003725836,0.0002890753,0.014773867,0.056994107,0.22529356,0.03181972,0.058016386,0.034964792,0.5227108],"study_design_scores_gemma":[0.0005493239,0.0010584175,0.0070830276,0.00026064867,0.0002479884,0.004427295,0.012292474,0.7586468,0.035643637,0.04031656,0.13927805,0.0001957603],"about_ca_topic_score_codex":0.006365096,"about_ca_topic_score_gemma":0.011448596,"teacher_disagreement_score":0.006365096,"about_ca_system_score_codex":0.0014558599,"about_ca_system_score_gemma":0.0010001542,"threshold_uncertainty_score":0.02719444},"labels":[],"label_agreement":null},{"id":"W7125599132","doi":"10.1109/cascon66301.2025.00112","title":"Enhancing Retrieval-Augmented Generation with Document Link Structure for Multi-Hop Web Question Answering","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"Brock University","keywords":"Hyperlink; Web modeling; Semantic Web; Context (archaeology); Data Web; Web page; Exploit; Web standards; Semantic search; Process (computing)","score_opus":0.021209173633728567,"score_gpt":0.2895249919053821,"score_spread":0.2683158182716535,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125599132","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039679557,0.0029048899,0.91274333,0.0008210028,0.00015992063,0.00048905617,0.0028506964,0.03669078,0.0036607364],"genre_scores_gemma":[0.2996477,0.00081577373,0.6797192,0.00060112216,0.00015204532,0.00038073098,0.012799958,0.0012716694,0.004611817],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986865,0.00050212047,0.000077616525,0.00038993033,0.00026754843,0.00007626963],"domain_scores_gemma":[0.9971697,0.0016887652,0.00013712964,0.000632604,0.00029053554,0.00008124531],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014340892,0.0011204911,0.00088470447,0.00296244,0.00065054273,0.0014742329,0.0016565843,0.0017500228,0.0037074434],"category_scores_gemma":[0.007409141,0.0004491026,0.0015405738,0.00214821,0.00065005996,0.0037108827,0.0021682428,0.0015530667,0.002695138],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000616332,0.0007646839,0.005319416,0.001271665,0.00023559392,0.000604894,0.0015153078,0.12259324,0.0384862,0.018375268,0.035047926,0.7751695],"study_design_scores_gemma":[0.00011691414,0.00020476828,0.0010734678,0.00004984784,0.00012216352,0.00035283505,0.00027830005,0.9219523,0.01736523,0.034515705,0.023909139,0.000059304042],"about_ca_topic_score_codex":0.0071539297,"about_ca_topic_score_gemma":0.01159033,"teacher_disagreement_score":0.0071539297,"about_ca_system_score_codex":0.00079221744,"about_ca_system_score_gemma":0.0010849604,"threshold_uncertainty_score":0.014224589},"labels":[],"label_agreement":null},{"id":"W7125620902","doi":"10.1109/cascon66301.2025.00037","title":"NL in the Middle: Code Translation with LLMs and Intermediate Representations","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Translation (biology); Syntax; Code (set theory); Abstract syntax; Natural language; Abstract syntax tree; Machine translation","score_opus":0.060226446811779556,"score_gpt":0.29567657845569906,"score_spread":0.2354501316439195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125620902","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10875738,0.0013541959,0.7412302,0.0021263002,0.00053093315,0.00034589588,0.0036556665,0.12944955,0.012549893],"genre_scores_gemma":[0.4616501,0.000404508,0.5126723,0.00070009957,0.000097573065,0.0002847679,0.010403372,0.008706788,0.0050805616],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99706143,0.0015226672,0.00020222412,0.0005518631,0.00048319934,0.00017868252],"domain_scores_gemma":[0.9873317,0.0066042207,0.0005633481,0.0037243667,0.0015000683,0.00027635338],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002970935,0.0016155929,0.0008041622,0.0013551835,0.0006928223,0.0027952604,0.0015878251,0.0017540974,0.0077085295],"category_scores_gemma":[0.029955696,0.00062536786,0.0009659873,0.0014903026,0.0009732271,0.005658234,0.0028336248,0.002583418,0.0046508964],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002173264,0.00070259505,0.0072908746,0.0020493788,0.00023653827,0.0011961659,0.0041920613,0.11013477,0.033175226,0.041730527,0.076106206,0.7210125],"study_design_scores_gemma":[0.00028741575,0.0005238977,0.0010699593,0.00021137061,0.00013026377,0.00045360238,0.00096679677,0.8313207,0.039046593,0.06726602,0.05861296,0.00011045815],"about_ca_topic_score_codex":0.004336568,"about_ca_topic_score_gemma":0.0062789046,"teacher_disagreement_score":0.0077085295,"about_ca_system_score_codex":0.0010261622,"about_ca_system_score_gemma":0.0019554705,"threshold_uncertainty_score":0.025787592},"labels":[],"label_agreement":null},{"id":"W7125701415","doi":"10.1109/medai67139.2025.00015","title":"Fine-Tuning Pre-trained Transformer-Based Models for Sentence-Level Medical Text Classification","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Dice; Class (philosophy); Transferability; Function (biology); Test (biology); Randomized controlled trial","score_opus":0.09921765004792034,"score_gpt":0.3202331665931665,"score_spread":0.22101551654524615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125701415","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15015996,0.0032988263,0.8176604,0.0012639646,0.00064999255,0.0006552186,0.0035254774,0.017660974,0.0051251603],"genre_scores_gemma":[0.69855887,0.0011279715,0.27856794,0.0009906456,0.0002810055,0.00076044543,0.011984109,0.0006851744,0.007043865],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99914265,0.00027416387,0.00008929415,0.00026385678,0.00014151439,0.00008844874],"domain_scores_gemma":[0.9979899,0.0010622131,0.00010147639,0.0001892568,0.000577327,0.00007977128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002472141,0.0013130852,0.00074423495,0.0012042829,0.00034120117,0.0009391271,0.001385468,0.0008561446,0.0023922815],"category_scores_gemma":[0.0078050024,0.00034305576,0.0012010427,0.00069612847,0.00036583966,0.0019352243,0.0011214874,0.0018750741,0.0027226303],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008631102,0.00042509587,0.0094747115,0.0005951787,0.00034718384,0.00041427746,0.0004064952,0.18347545,0.045167625,0.0034972357,0.025556775,0.72977686],"study_design_scores_gemma":[0.00004700509,0.00025791384,0.0015300999,0.000049826325,0.00011145225,0.00020987993,0.000099034136,0.96718264,0.02161208,0.0039015273,0.0049637826,0.000034798413],"about_ca_topic_score_codex":0.0044321236,"about_ca_topic_score_gemma":0.0077512045,"teacher_disagreement_score":0.0044321236,"about_ca_system_score_codex":0.00091357075,"about_ca_system_score_gemma":0.0015560167,"threshold_uncertainty_score":0.0130741},"labels":[],"label_agreement":null},{"id":"W7125777035","doi":"10.21428/594757db.1b7841bc","title":"Aligning Language Models Using Multi-Objective Deep Reinforcement Learning","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University; Brock University","funders":"","keywords":"Helpfulness; Reinforcement learning; Natural language; Language model; SIGNAL (programming language); Reinforcement; Language acquisition","score_opus":0.030782286661813492,"score_gpt":0.2911095375561579,"score_spread":0.26032725089434444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125777035","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05191633,0.0003941493,0.94357264,0.00028119757,0.000078022844,0.00006919274,0.00005918222,0.0016256932,0.002003568],"genre_scores_gemma":[0.8778431,0.00013994938,0.11884384,0.0002584012,0.000046620014,0.00013828462,0.00019145968,0.00019177089,0.0023466465],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993345,0.0002633096,0.000029492581,0.00018632953,0.00011123444,0.00007496955],"domain_scores_gemma":[0.99841046,0.0010448083,0.00017887908,0.00010282855,0.00017519371,0.00008776154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013378405,0.001175326,0.00089445926,0.00045905402,0.00030218903,0.00085034897,0.0011972067,0.001030194,0.0018581757],"category_scores_gemma":[0.0047098584,0.0005332411,0.0005956707,0.00036565817,0.0006708129,0.001510425,0.0011540129,0.0021765647,0.00057410204],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007549951,0.00010278407,0.0006941979,0.000047707996,0.000049687536,0.00006114551,0.00008363661,0.9324896,0.0029747204,0.0034353265,0.00095217745,0.05903343],"study_design_scores_gemma":[0.0000049976784,0.000013186668,0.00002241772,0.0000015635733,0.0000024253152,0.0000036028966,0.0000030419976,0.99827075,0.00028013214,0.0013187964,0.00007679069,0.00000223432],"about_ca_topic_score_codex":0.0043868953,"about_ca_topic_score_gemma":0.0058477484,"teacher_disagreement_score":0.0043868953,"about_ca_system_score_codex":0.0009419889,"about_ca_system_score_gemma":0.0010079824,"threshold_uncertainty_score":0.0087227225},"labels":[],"label_agreement":null},{"id":"W7125920169","doi":"10.21428/594757db.5968b33e","title":"MAC: Multimodal Attentive Contrastive Learning Framework","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Multimodal learning; Feature learning; Mechanism (biology); Joint (building); Multimodality; Representation (politics); Contrast (vision)","score_opus":0.008993020724235603,"score_gpt":0.26377009045569294,"score_spread":0.25477706973145736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125920169","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010971175,0.00038164516,0.98546416,0.00030803517,0.000044862492,0.0000677211,0.00014327644,0.001090585,0.0015284956],"genre_scores_gemma":[0.5785434,0.000532504,0.41124573,0.0011307442,0.0002722836,0.0005769051,0.0009868952,0.00035442843,0.006357001],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921227,0.00030878742,0.000024973511,0.00025771846,0.00012166192,0.00007462116],"domain_scores_gemma":[0.99867237,0.0007809661,0.00009246034,0.00015425737,0.00022219292,0.00007776762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021219815,0.001466339,0.00085430243,0.0009841819,0.00047188456,0.0010239839,0.00243987,0.0015781643,0.0034451901],"category_scores_gemma":[0.005465686,0.00049114827,0.001035252,0.00061113725,0.000986005,0.0020368176,0.0026358939,0.0027060066,0.00104338],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004794666,0.00038091018,0.0028642232,0.00026543395,0.00026201416,0.00027983487,0.00042654524,0.3918006,0.019187437,0.037635464,0.010980821,0.5354372],"study_design_scores_gemma":[0.000014976926,0.000083953004,0.00022257924,0.000014894839,0.000023422257,0.000044048695,0.000019794295,0.97945774,0.0022168623,0.01652593,0.0013632955,0.000012503134],"about_ca_topic_score_codex":0.0024295235,"about_ca_topic_score_gemma":0.0034694388,"teacher_disagreement_score":0.0034451901,"about_ca_system_score_codex":0.00087706244,"about_ca_system_score_gemma":0.00075472164,"threshold_uncertainty_score":0.011525273},"labels":[],"label_agreement":null},{"id":"W7125974680","doi":"10.1109/smc58881.2025.11343466","title":"Towards Robust Retrieval-Augmented Generation Based on Knowledge Graph: A Comparative Analysis","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cégep de l'Outaouais; Institut National de la Recherche Scientifique","funders":"","keywords":"Robustness (evolution); Counterfactual thinking; Testbed; Personalization; Natural language generation; RGB color model","score_opus":0.08601229894000449,"score_gpt":0.32554630638258597,"score_spread":0.23953400744258146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7125974680","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50751346,0.011523323,0.38543805,0.0014178661,0.00054673565,0.001646059,0.0056755715,0.058664802,0.027574165],"genre_scores_gemma":[0.7690032,0.0012544648,0.20890553,0.00051235396,0.00008262623,0.0003666351,0.013932611,0.0021016398,0.003840962],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98887,0.0053046136,0.0006941384,0.0017732952,0.0028319682,0.00052598154],"domain_scores_gemma":[0.97293764,0.0161985,0.00084205903,0.00697296,0.0026144013,0.00043451373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009972892,0.0017685404,0.0014033922,0.0029996105,0.00065002555,0.0029327408,0.002759572,0.0024429276,0.0039952486],"category_scores_gemma":[0.037704833,0.0005081495,0.001154429,0.0019245645,0.0015207995,0.004908323,0.0033239394,0.0017119577,0.002075569],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039546574,0.0010077497,0.008574057,0.0046362476,0.00090140046,0.0009384781,0.0014894628,0.30349073,0.037968848,0.012508407,0.021775981,0.60275406],"study_design_scores_gemma":[0.00037978767,0.0018933516,0.006650452,0.00020491164,0.0004583222,0.0007151396,0.0009092049,0.9016187,0.05055,0.013363211,0.023082105,0.00017478848],"about_ca_topic_score_codex":0.009670102,"about_ca_topic_score_gemma":0.0070383893,"teacher_disagreement_score":0.009972892,"about_ca_system_score_codex":0.0014685434,"about_ca_system_score_gemma":0.0015818585,"threshold_uncertainty_score":0.052742302},"labels":[],"label_agreement":null},{"id":"W7126370870","doi":"10.21428/594757db.2a2b97c3","title":"Triple-Aware Reasoning: A Retrieval-Augmented GenerationApproach for Enhancing Question-Answering Tasks withKnowledge Graphs and Large Language Models","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Language model; Knowledge graph; Inference; Relevance (law); Question answering; Universal Networking Language; Data modeling; Natural language; Graph","score_opus":0.02018945031357532,"score_gpt":0.2802207549337358,"score_spread":0.2600313046201605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126370870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059053763,0.00039888886,0.9856908,0.0002990792,0.00005015457,0.00019900904,0.00043137118,0.0060161757,0.0010091835],"genre_scores_gemma":[0.14228621,0.0004399968,0.84942186,0.00054571877,0.000097779026,0.00033323572,0.0036796085,0.0005291995,0.002666403],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747145,0.00092655316,0.00014992592,0.0007343317,0.00058301125,0.00013478345],"domain_scores_gemma":[0.9968555,0.001585431,0.000165698,0.0009394461,0.00035408756,0.000099833946],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00252566,0.0022598903,0.0013329929,0.0031157413,0.00086175953,0.0012987439,0.0036693525,0.002051826,0.004236805],"category_scores_gemma":[0.006406968,0.000835878,0.0033361793,0.002370253,0.0009859566,0.0048125265,0.0030354187,0.0025852348,0.0020359145],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003562603,0.00085437065,0.0020269235,0.00078736164,0.0003327597,0.0008307612,0.0010737599,0.09183082,0.033596355,0.027151987,0.026987256,0.8141714],"study_design_scores_gemma":[0.00008087369,0.00010093586,0.00042137667,0.0000388145,0.00014078873,0.00028868916,0.00013092488,0.9264774,0.0112735545,0.05106119,0.009917285,0.00006824082],"about_ca_topic_score_codex":0.0065410375,"about_ca_topic_score_gemma":0.011701432,"teacher_disagreement_score":0.0065410375,"about_ca_system_score_codex":0.00093061617,"about_ca_system_score_gemma":0.0017160249,"threshold_uncertainty_score":0.014173567},"labels":[],"label_agreement":null},{"id":"W7126381417","doi":"10.21428/594757db.3f69a234","title":"Advancements in Large Language Models Through Employment of Retrieval Augmented Generation and its Enhancements for Assistance in Academia","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Popularity; Domain (mathematical analysis); Language model; Data modeling; Topic model","score_opus":0.0412717576777256,"score_gpt":0.34182395355652867,"score_spread":0.30055219587880305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126381417","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014721928,0.0013745663,0.93940765,0.0016180696,0.0002543107,0.00042295334,0.00084895716,0.03683167,0.0045198635],"genre_scores_gemma":[0.21751614,0.0011464726,0.7683095,0.0011632,0.0002939711,0.00046812385,0.0022154253,0.0016471902,0.007240028],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9939142,0.0036482392,0.00041235777,0.0009945625,0.00084522733,0.00018545998],"domain_scores_gemma":[0.9804904,0.010874276,0.0009475802,0.0054462743,0.0018552958,0.00038625082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009368782,0.0012397812,0.0010383991,0.0016395727,0.000571472,0.003134879,0.002382916,0.0015951298,0.007868245],"category_scores_gemma":[0.023850081,0.00064610754,0.0017015728,0.001375515,0.0009989114,0.005569933,0.0026028696,0.0026328608,0.008358003],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091676996,0.00047724336,0.004325201,0.0013097728,0.00020473683,0.00034185752,0.0022243215,0.021015175,0.045887303,0.02402647,0.02282074,0.87645036],"study_design_scores_gemma":[0.0002681811,0.0013361764,0.0034394942,0.0003580052,0.00043588804,0.0013681793,0.0008931695,0.6934787,0.084416814,0.04951411,0.16403571,0.00045559157],"about_ca_topic_score_codex":0.002351702,"about_ca_topic_score_gemma":0.0026640275,"teacher_disagreement_score":0.009368782,"about_ca_system_score_codex":0.0009903716,"about_ca_system_score_gemma":0.0016469457,"threshold_uncertainty_score":0.049547434},"labels":[],"label_agreement":null},{"id":"W7126402336","doi":"10.18653/v1/2024.findings-eacl.43","title":"Capturing the Relationship Between Sentence Triplets for LLM and Human-Generated Texts to Enhance Sentence Embeddings","year":2024,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"Institute for Information and Communications Technology Promotion; Ministry of Science and ICT, South Korea; Korea Advanced Institute of Science and Technology","keywords":"Sentence; Feature (linguistics); Semantics (computer science); Term (time)","score_opus":0.07427577997439558,"score_gpt":0.3428753693139701,"score_spread":0.2685995893395745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126402336","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18930729,0.0015412793,0.7676258,0.0009185699,0.00086024334,0.0005168585,0.008903867,0.024437016,0.0058890632],"genre_scores_gemma":[0.48063928,0.00036683452,0.4804675,0.00042821298,0.0001759509,0.00073685846,0.029368293,0.0012039909,0.0066130203],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985184,0.00058445,0.00009188469,0.0005576906,0.00018073192,0.00006684569],"domain_scores_gemma":[0.99681646,0.0014993611,0.00025086332,0.0008821159,0.00043357696,0.00011759656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024942837,0.0015358924,0.0005736076,0.0011630318,0.00042732194,0.0009784801,0.0010726409,0.001112102,0.0059352485],"category_scores_gemma":[0.011678153,0.00033538672,0.00095278025,0.00095593045,0.00054501346,0.0024396772,0.0017969897,0.001824441,0.004924945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000815044,0.00058266683,0.0154225165,0.0011516435,0.00025935922,0.0004601965,0.0009059003,0.055557486,0.072807245,0.007451516,0.050840788,0.79374564],"study_design_scores_gemma":[0.000062811334,0.0005197154,0.0064079813,0.00011132402,0.0000734119,0.0005607349,0.0003773902,0.91863143,0.041543733,0.010172493,0.02147248,0.00006641916],"about_ca_topic_score_codex":0.0014141707,"about_ca_topic_score_gemma":0.003701799,"teacher_disagreement_score":0.0059352485,"about_ca_system_score_codex":0.0006127339,"about_ca_system_score_gemma":0.000603678,"threshold_uncertainty_score":0.01985538},"labels":[],"label_agreement":null},{"id":"W7126404497","doi":"10.21428/594757db.834c24c6","title":"Intra-Layer Recurrence in Transformers for Language Modeling","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Algoma University","funders":"","keywords":"Transformer; Language model; Natural language; Modeling language; Current transformer","score_opus":0.03016604409396594,"score_gpt":0.3008865339034704,"score_spread":0.27072048980950447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126404497","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016274933,0.00017362628,0.9802957,0.00009363327,0.000014714877,0.000034223194,0.00007357186,0.0019760001,0.0010635656],"genre_scores_gemma":[0.5735059,0.00040115067,0.42055082,0.00015280103,0.00004142077,0.00013298608,0.00043143908,0.0009031349,0.0038803439],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99935097,0.00024005986,0.000051014882,0.00015161972,0.00013062163,0.00007584562],"domain_scores_gemma":[0.99804556,0.0011438222,0.00017037985,0.00040415596,0.00017739594,0.000058701273],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001132735,0.0007520331,0.0005993299,0.0005571638,0.0004224795,0.0011801478,0.0012962216,0.0005968477,0.0041616075],"category_scores_gemma":[0.0057496494,0.0005353443,0.0011006393,0.0005611917,0.000796446,0.003920056,0.0011084435,0.0014251779,0.0012932103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037569948,0.00011853156,0.0024750417,0.00036618495,0.00014419945,0.00034958645,0.00077366404,0.49430642,0.0397788,0.14874998,0.0042370213,0.30832493],"study_design_scores_gemma":[0.000010744175,0.00006178791,0.00009670552,0.000011773135,0.000034534627,0.00008439656,0.000038875827,0.93458784,0.011665424,0.050622266,0.0027716216,0.000014081691],"about_ca_topic_score_codex":0.0037273208,"about_ca_topic_score_gemma":0.009263329,"teacher_disagreement_score":0.0041616075,"about_ca_system_score_codex":0.00091528846,"about_ca_system_score_gemma":0.0010087341,"threshold_uncertainty_score":0.0139219165},"labels":[],"label_agreement":null},{"id":"W7126406861","doi":"10.21428/594757db.c4952d65","title":"Discovering and Evaluating Politeness Editing Strategies","year":2025,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; Toronto Metropolitan University","funders":"","keywords":"Politeness; Rewriting; Natural (archaeology); Natural language; Phenomenon; Natural language generation","score_opus":0.03139444771489772,"score_gpt":0.33202797958433133,"score_spread":0.3006335318694336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126406861","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94202,0.0012247124,0.046251085,0.00026224653,0.000067321445,0.00024839965,0.0011549944,0.0023445666,0.0064268145],"genre_scores_gemma":[0.9592757,0.00029554783,0.034442306,0.000039821487,0.00003483974,0.00012532879,0.0040893974,0.00014237901,0.0015546401],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955296,0.0023231274,0.00030288316,0.00096620497,0.0006501781,0.00022804983],"domain_scores_gemma":[0.97557676,0.018625444,0.0015875945,0.0014487742,0.0021292006,0.0006322065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004532092,0.00081513607,0.0006758876,0.0043260865,0.00057844364,0.003040362,0.0005874646,0.0010301145,0.0011759093],"category_scores_gemma":[0.028768202,0.00034989562,0.0006702327,0.0016803186,0.0005735218,0.002481698,0.0012673343,0.0010862538,0.0010385999],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011726558,0.00088283396,0.24698019,0.0017863226,0.0005055224,0.00044921463,0.00964858,0.024832197,0.044073347,0.0034337565,0.0077730594,0.6584623],"study_design_scores_gemma":[0.00012035135,0.001095786,0.23189701,0.0001935766,0.00038043046,0.000871701,0.006414951,0.6959262,0.041230723,0.0050492226,0.016635178,0.00018492484],"about_ca_topic_score_codex":0.0038370923,"about_ca_topic_score_gemma":0.005329154,"teacher_disagreement_score":0.004532092,"about_ca_system_score_codex":0.0008520485,"about_ca_system_score_gemma":0.0007117955,"threshold_uncertainty_score":0.023968339},"labels":[],"label_agreement":null},{"id":"W7126410033","doi":"10.21428/594757db.1c863f43","title":"Large Language Models are Incoherent Storytellers","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Variety (cybernetics); Narrative; Natural language; Natural language generation; Natural (archaeology); Language planning; Range (aeronautics)","score_opus":0.024723074066221155,"score_gpt":0.2581021526382649,"score_spread":0.23337907857204376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126410033","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030860443,0.00081158907,0.94066703,0.0050451164,0.00011599832,0.00014509779,0.00062308315,0.0020803942,0.019651234],"genre_scores_gemma":[0.5382479,0.0010320244,0.4341062,0.0016734931,0.00030734562,0.00075476477,0.0021390007,0.0014086474,0.020330679],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99449426,0.0028730065,0.00031499699,0.001003724,0.0011560789,0.000157859],"domain_scores_gemma":[0.9759311,0.016995067,0.0011617346,0.004162789,0.0012867413,0.00046249875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052305507,0.000959516,0.00078940473,0.001502626,0.0011788921,0.0055084853,0.0021378968,0.0019190045,0.0065948875],"category_scores_gemma":[0.03968281,0.0014477796,0.0009514495,0.0007912962,0.004152685,0.012429848,0.0048104143,0.0034682571,0.002042684],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000097041935,0.00006337203,0.0016459618,0.00031379252,0.00010078262,0.00045079913,0.004210213,0.032102548,0.0031857768,0.9163911,0.0065672384,0.034871325],"study_design_scores_gemma":[0.000027409138,0.000031075095,0.00017567119,0.00005409105,0.000039424125,0.00018966761,0.00042781202,0.12337939,0.0015010564,0.8507382,0.023409382,0.000026906439],"about_ca_topic_score_codex":0.001165883,"about_ca_topic_score_gemma":0.0020535686,"teacher_disagreement_score":0.0065948875,"about_ca_system_score_codex":0.0011567654,"about_ca_system_score_gemma":0.0009113841,"threshold_uncertainty_score":0.027662098},"labels":[],"label_agreement":null},{"id":"W7126411529","doi":"10.18653/v1/2024.codi-1.11","title":"Exploring Soft-Label Training for Implicit Discourse Relation Recognition","year":2024,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Training (meteorology); Relation (database); Feature (linguistics); Action (physics); Interpretation (philosophy)","score_opus":0.38760953990793007,"score_gpt":0.3513438588791463,"score_spread":0.03626568102878375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126411529","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05844485,0.0016623567,0.9274226,0.0011354097,0.00032418958,0.00017175374,0.000996066,0.0070218476,0.0028209034],"genre_scores_gemma":[0.58241177,0.0006065732,0.4000576,0.00061591476,0.0004986322,0.00039831942,0.0058363797,0.0010481115,0.008526701],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967019,0.0015231861,0.00020919806,0.00088428403,0.0003980373,0.000283301],"domain_scores_gemma":[0.9794632,0.016970366,0.00043889915,0.0016066564,0.0011551934,0.00036554845],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004226437,0.0017468135,0.0016826534,0.0017453624,0.0013723196,0.0029422,0.0034402497,0.0029243047,0.0053621586],"category_scores_gemma":[0.016789654,0.0008696495,0.0014237392,0.0017943486,0.0009992641,0.006101583,0.0035493593,0.0057363897,0.0032973534],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011256335,0.00073721784,0.00364683,0.0006932307,0.00019781734,0.000253629,0.0016357343,0.05474496,0.02178664,0.016811341,0.015917605,0.88244927],"study_design_scores_gemma":[0.000042821786,0.00008744789,0.00040663048,0.00006592941,0.000048246904,0.000043915312,0.00021743732,0.96756685,0.0067453957,0.021578368,0.0031755061,0.000021545891],"about_ca_topic_score_codex":0.0067008757,"about_ca_topic_score_gemma":0.013302203,"teacher_disagreement_score":0.0067008757,"about_ca_system_score_codex":0.0010865415,"about_ca_system_score_gemma":0.0022116664,"threshold_uncertainty_score":0.022351801},"labels":[],"label_agreement":null},{"id":"W7126425016","doi":"10.18653/v1/2024.findings-eacl.128","title":"Jigsaw Pieces of Meaning: Modeling Discourse Coherence with Informed Negative Sample Synthesis","year":2024,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Coherence (philosophical gambling strategy); Jigsaw; Sample (material); Set (abstract data type)","score_opus":0.04639313863466781,"score_gpt":0.290176069567851,"score_spread":0.2437829309331832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126425016","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.097895026,0.00057871785,0.89605683,0.0004583563,0.00007864863,0.000099769764,0.00027862494,0.0015336829,0.0030203194],"genre_scores_gemma":[0.77791643,0.00016061119,0.2168207,0.00021786035,0.000064230786,0.0002019673,0.0008600355,0.00051118026,0.0032470508],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99900985,0.00047837768,0.000036935908,0.00028965945,0.00012997708,0.000055248263],"domain_scores_gemma":[0.9964077,0.0025331706,0.00020585919,0.00041880447,0.00033378703,0.00010074166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021607627,0.0009751892,0.0006752477,0.0007047823,0.0006393441,0.0014722605,0.0012986356,0.0010141032,0.0026787657],"category_scores_gemma":[0.012481451,0.00052456034,0.00076708436,0.00047793132,0.0014265432,0.0028915657,0.0019153274,0.0018667928,0.00068050507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011561833,0.00029227667,0.006293648,0.00053771416,0.00020899998,0.00049194176,0.0041959235,0.44080058,0.041359983,0.06895843,0.0062095555,0.4294947],"study_design_scores_gemma":[0.00004138792,0.00008413364,0.00055237155,0.000026863554,0.000033495937,0.000061418614,0.00018180128,0.94735205,0.006313117,0.042753905,0.0025800848,0.000019308136],"about_ca_topic_score_codex":0.002619457,"about_ca_topic_score_gemma":0.004929214,"teacher_disagreement_score":0.0026787657,"about_ca_system_score_codex":0.00078072114,"about_ca_system_score_gemma":0.00074488844,"threshold_uncertainty_score":0.011427343},"labels":[],"label_agreement":null},{"id":"W7126443791","doi":"10.21428/594757db.a9152d5a","title":"SoftSkillQG: Robust Skill-Targeted Reading ComprehensionQuestion Generation Using Soft-prompts","year":2024,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Variety (cybernetics); Context (archaeology); Quality (philosophy); Comprehension; Reading (process); Reading comprehension; Language understanding","score_opus":0.06561443085860268,"score_gpt":0.28378855641263656,"score_spread":0.21817412555403387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7126443791","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08714774,0.002913245,0.48241732,0.0009995055,0.0007496666,0.0021843952,0.026911873,0.38879144,0.0078849215],"genre_scores_gemma":[0.3081263,0.000769907,0.5612859,0.001199408,0.00022851514,0.0028254853,0.10503431,0.008660047,0.011870104],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968183,0.0012894226,0.0002247843,0.0010090951,0.00051876577,0.00013952823],"domain_scores_gemma":[0.99325037,0.0037423293,0.000339624,0.0012799723,0.0011019728,0.00028578186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002835024,0.0028451073,0.0012754379,0.0017308252,0.00042004092,0.0016255536,0.0028128296,0.0019555422,0.010354432],"category_scores_gemma":[0.012340313,0.00059872755,0.0012908151,0.0009341603,0.00063599204,0.002780896,0.0030182707,0.002166569,0.009870418],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010594383,0.00094113126,0.0061485665,0.0035377345,0.00020693446,0.00058896764,0.0022028482,0.019208368,0.05282081,0.0034557958,0.1795547,0.7302747],"study_design_scores_gemma":[0.0015974718,0.0017824798,0.010084759,0.00032493658,0.00021434047,0.0011034486,0.0015367775,0.66382945,0.109999314,0.021067902,0.18810412,0.0003550968],"about_ca_topic_score_codex":0.002622783,"about_ca_topic_score_gemma":0.0048801024,"teacher_disagreement_score":0.010354432,"about_ca_system_score_codex":0.00083834806,"about_ca_system_score_gemma":0.0012503732,"threshold_uncertainty_score":0.034639},"labels":[],"label_agreement":null},{"id":"W7128191878","doi":"10.1145/3745764","title":"NLPerturbator: Studying the Robustness of Code LLMs to Natural Language Variations","year":2025,"lang":"en","type":"article","venue":"ACM Transactions on Software Engineering and Methodology","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"National Key Research and Development Program of China","keywords":"Robustness (evolution); Natural language; Coding (social sciences); Code (set theory); Natural language generation","score_opus":0.056644086971520985,"score_gpt":0.31641526681795956,"score_spread":0.2597711798464386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7128191878","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.630388,0.0017352307,0.34624124,0.0006959021,0.00019628258,0.0008765706,0.0014894708,0.01426695,0.0041104085],"genre_scores_gemma":[0.8875883,0.00031626556,0.1068349,0.00028931157,0.000052544543,0.00058630196,0.002169569,0.0013015383,0.00086123013],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9795815,0.009436825,0.0013828434,0.0037204975,0.0052235667,0.00065472384],"domain_scores_gemma":[0.8282169,0.12024109,0.015274989,0.02738224,0.0076454272,0.0012393644],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013466648,0.0014664375,0.00082554814,0.0027865772,0.00081969576,0.0020695482,0.0022735104,0.0015791247,0.0012186419],"category_scores_gemma":[0.1662024,0.0008184973,0.0010063843,0.0016611954,0.0027639002,0.0047285375,0.003065491,0.0025697716,0.00078540656],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003193958,0.0012084165,0.08667507,0.0026194712,0.0009984118,0.0007628776,0.004749591,0.44120932,0.08713088,0.012605271,0.008939668,0.34990695],"study_design_scores_gemma":[0.00014943589,0.0015572661,0.023015615,0.00012308448,0.00022952682,0.0005570541,0.0008071381,0.8969845,0.053193145,0.015900504,0.0072917505,0.00019100458],"about_ca_topic_score_codex":0.0050553796,"about_ca_topic_score_gemma":0.0033921795,"teacher_disagreement_score":0.98653334,"about_ca_system_score_codex":0.0019368596,"about_ca_system_score_gemma":0.0018936384,"threshold_uncertainty_score":0.071219265},"labels":[],"label_agreement":null},{"id":"W7128810986","doi":"10.26615/978-954-452-106-6-011","title":"Exploring Language in Different Daily Time Segments Through Text Prediction and Language Modeling","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"St. Francis Xavier University","keywords":"Language model; Natural language; Modeling language; Language identification; Feature (linguistics); Data modeling","score_opus":0.05274004071677961,"score_gpt":0.2675619732917397,"score_spread":0.21482193257496007,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7128810986","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7829142,0.0019686033,0.1980166,0.0015207859,0.00024985868,0.00019159449,0.005797785,0.0045217206,0.004818888],"genre_scores_gemma":[0.94399804,0.0006063931,0.047829043,0.00015823015,0.00013042199,0.00013088001,0.004465803,0.0001499308,0.0025313443],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99941766,0.00019901281,0.00004376541,0.00020974506,0.00007468933,0.000055139466],"domain_scores_gemma":[0.9966137,0.002498221,0.00029722272,0.00018605629,0.00027850005,0.00012628411],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011423413,0.0011400599,0.00051451137,0.0018859602,0.000365101,0.0011832865,0.00077895774,0.00074306916,0.0013252445],"category_scores_gemma":[0.004955878,0.00030090538,0.0011348607,0.0016176676,0.00024747508,0.0025347942,0.00051603763,0.0012387619,0.001627507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015828311,0.0011634263,0.14739007,0.0008143288,0.00060734036,0.0010727573,0.0026977733,0.25268796,0.035149973,0.0036027625,0.01728158,0.53594923],"study_design_scores_gemma":[0.00001608575,0.00012965914,0.0108250575,0.000033142016,0.00007084612,0.00015408978,0.0003250067,0.980361,0.0029781386,0.0024744063,0.0025971243,0.00003537531],"about_ca_topic_score_codex":0.011320983,"about_ca_topic_score_gemma":0.018047554,"teacher_disagreement_score":0.011320983,"about_ca_system_score_codex":0.00054960436,"about_ca_system_score_gemma":0.00072202366,"threshold_uncertainty_score":0.022510171},"labels":[],"label_agreement":null},{"id":"W7130555830","doi":"10.1109/fllm67465.2025.11391225","title":"Emotion-Aware Prompting for Large Language Models","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Artificial Intelligence in Medicine (Canada); Université TÉLUQ","funders":"","keywords":"Context (archaeology); Natural language; Language model; Natural (archaeology); Benchmark (surveying); Mental model; Context model; Natural language understanding","score_opus":0.029708949427258897,"score_gpt":0.2969832034087523,"score_spread":0.2672742539814934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7130555830","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018248674,0.0005292333,0.9473225,0.00055497035,0.00012435167,0.00016528607,0.0013759786,0.03092604,0.00075295795],"genre_scores_gemma":[0.37860927,0.00043687937,0.6097933,0.000493396,0.00019388949,0.00070707203,0.005663855,0.0015108662,0.0025915322],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99847454,0.00085165363,0.00009288804,0.00035874787,0.00015803364,0.00006412848],"domain_scores_gemma":[0.99458534,0.0039835838,0.00024123919,0.00072771613,0.00033762096,0.00012441695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024819614,0.0013294922,0.00078565767,0.00056481466,0.00040947302,0.0011264092,0.0011824584,0.0009862756,0.004092133],"category_scores_gemma":[0.01267086,0.00046144178,0.0011826268,0.00047561215,0.00041153393,0.0020087326,0.0015665552,0.0022558277,0.0025700366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013170002,0.0004563486,0.0039019438,0.0010131293,0.00020497943,0.000412312,0.0011634869,0.19443539,0.058113296,0.016265001,0.03701187,0.6857053],"study_design_scores_gemma":[0.00007650958,0.00012318116,0.000471653,0.000021704502,0.000035699733,0.00007386976,0.00011956835,0.9493475,0.010644102,0.030761406,0.008294265,0.000030453006],"about_ca_topic_score_codex":0.001676711,"about_ca_topic_score_gemma":0.0032433032,"teacher_disagreement_score":0.004092133,"about_ca_system_score_codex":0.0007038795,"about_ca_system_score_gemma":0.0010072485,"threshold_uncertainty_score":0.013689518},"labels":[],"label_agreement":null},{"id":"W7130611624","doi":"10.1109/fllm67465.2025.11391199","title":"Enhancing the Reliability of Large Language Models in Specialized Domains by Balancing Internal and External Knowledge","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Reliability (semiconductor); Focus (optics); Measure (data warehouse); Language model; Training (meteorology); Training set","score_opus":0.00894258193979811,"score_gpt":0.2789350916978804,"score_spread":0.26999250975808226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7130611624","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05430211,0.003615438,0.9284152,0.0008154659,0.0001491361,0.00022653808,0.0004187916,0.00954047,0.0025168771],"genre_scores_gemma":[0.6664761,0.0013963555,0.3222165,0.0010155651,0.0004278644,0.00032787985,0.0029732492,0.0012134698,0.003953076],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99711525,0.0012450246,0.00019919996,0.00081787404,0.00044103875,0.00018158235],"domain_scores_gemma":[0.98361397,0.011777021,0.0007200992,0.0024525658,0.001077555,0.0003588556],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00585818,0.0028852306,0.0024724985,0.0024157418,0.0009802851,0.0025589387,0.0036267997,0.0028281196,0.0020887055],"category_scores_gemma":[0.026799684,0.0010918524,0.0019063026,0.0014610717,0.0013533237,0.0048877755,0.0033610002,0.003630689,0.002580218],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009363004,0.00043447438,0.005302254,0.00090007234,0.0005298572,0.0006212128,0.00089008134,0.29505926,0.02104017,0.005393347,0.013157743,0.6557353],"study_design_scores_gemma":[0.000054327196,0.00008969032,0.00050093175,0.00004016195,0.000082147715,0.00017050435,0.00010777023,0.98423785,0.005178542,0.008015716,0.0014903685,0.000031947424],"about_ca_topic_score_codex":0.0043344502,"about_ca_topic_score_gemma":0.006679205,"teacher_disagreement_score":0.00585818,"about_ca_system_score_codex":0.0011687841,"about_ca_system_score_gemma":0.0017969034,"threshold_uncertainty_score":0.030981421},"labels":[],"label_agreement":null},{"id":"W7131116583","doi":"10.1109/icairc68035.2025.11384878","title":"Two-Stage Fine-Tuning for Retrieval-Augmented Large Language Models via Retriever and Prompt Optimization","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Language model; Generalization; Generator (circuit theory); Labrador Retriever; Protocol (science); Set (abstract data type)","score_opus":0.02058260023157539,"score_gpt":0.2826260557435897,"score_spread":0.2620434555120143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7131116583","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025042133,0.0010739148,0.9176107,0.00047407814,0.0002505706,0.00039570953,0.0005636218,0.051163886,0.0034253353],"genre_scores_gemma":[0.40689915,0.0005067031,0.5716441,0.0015144979,0.00026042602,0.0011036168,0.0037408082,0.0061191744,0.008211453],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99819654,0.0007424817,0.00010722123,0.0005466051,0.00025644156,0.00015080151],"domain_scores_gemma":[0.9970732,0.0017652166,0.00010437492,0.00064383756,0.00029948342,0.00011383022],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028206513,0.0021335466,0.0015822991,0.0010307943,0.00051445543,0.0016544525,0.0027959424,0.0018684524,0.0112438295],"category_scores_gemma":[0.010786053,0.000833737,0.0016958453,0.00071375386,0.0009733638,0.0027702686,0.0030354545,0.0033836227,0.008090205],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00097458955,0.0006073444,0.0028320951,0.0008824416,0.0002947176,0.00062734075,0.00063040777,0.2010508,0.049197063,0.011299315,0.037962265,0.69364154],"study_design_scores_gemma":[0.0002517213,0.00024089727,0.000501582,0.00004220303,0.000094348696,0.00029857774,0.000104975894,0.95914584,0.017298335,0.013544781,0.00841693,0.000059811748],"about_ca_topic_score_codex":0.0021735448,"about_ca_topic_score_gemma":0.0049272007,"teacher_disagreement_score":0.0112438295,"about_ca_system_score_codex":0.00071608427,"about_ca_system_score_gemma":0.0013654219,"threshold_uncertainty_score":0.037614346},"labels":[],"label_agreement":null},{"id":"W7131634265","doi":"10.1109/ickg66886.2025.00020","title":"Evaluating Knowledge Graph-Enhanced Context for Multiple-Choice Question Answering","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brock University","funders":"","keywords":"Relevance (law); Question answering; Selection (genetic algorithm); Context (archaeology); Scope (computer science); Key (lock); Knowledge representation and reasoning; Domain knowledge; General knowledge","score_opus":0.08508419509513357,"score_gpt":0.3952205903199808,"score_spread":0.31013639522484726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7131634265","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69785917,0.015536886,0.231937,0.0012479582,0.00063587615,0.0011156739,0.006154031,0.03740377,0.008109644],"genre_scores_gemma":[0.8574164,0.0011577035,0.13063452,0.0004899202,0.00011848367,0.00033280306,0.008253684,0.00020774094,0.0013887865],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978282,0.0009773526,0.00011007367,0.0006368856,0.00033637247,0.000111156],"domain_scores_gemma":[0.9915235,0.0067738937,0.00031715733,0.000565919,0.0004819833,0.0003375044],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024855435,0.0017272613,0.00094715605,0.0016554033,0.00041501425,0.0012965229,0.0018162177,0.0023239495,0.0036185947],"category_scores_gemma":[0.018622085,0.00030844184,0.001064954,0.000866352,0.0004601436,0.0042079487,0.0023131862,0.0017265773,0.0013785056],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0048322375,0.002121102,0.017351385,0.0037675393,0.0005562155,0.000564983,0.0016073163,0.10911665,0.023453526,0.0044208933,0.016463706,0.8157444],"study_design_scores_gemma":[0.00045322493,0.0028493323,0.008249071,0.00026042506,0.00038575567,0.0004925531,0.00080945046,0.942772,0.01435614,0.017730346,0.01152519,0.00011647052],"about_ca_topic_score_codex":0.0059748706,"about_ca_topic_score_gemma":0.007305108,"teacher_disagreement_score":0.0059748706,"about_ca_system_score_codex":0.000848206,"about_ca_system_score_gemma":0.0011480352,"threshold_uncertainty_score":0.01314497},"labels":[],"label_agreement":null},{"id":"W7132553631","doi":"","title":"Challenges in technical regulatory text variation detection","year":2025,"lang":"en","type":"article","venue":"NPARC","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Task (project management); Sentence; Representation (politics); Variation (astronomy); Natural language","score_opus":0.030606485729304217,"score_gpt":0.25811262500236765,"score_spread":0.22750613927306343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7132553631","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15614563,0.0036079776,0.8089484,0.011218781,0.00059584057,0.00033811035,0.0033918822,0.0076386095,0.00811478],"genre_scores_gemma":[0.5616702,0.0015131007,0.42240697,0.0014837688,0.00057439855,0.0003442941,0.008115025,0.0010183853,0.0028739262],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9860058,0.008307836,0.0008718406,0.002520178,0.0019744888,0.00031991795],"domain_scores_gemma":[0.9449271,0.042597722,0.002692836,0.00391274,0.0054543274,0.00041522697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012217549,0.0010035349,0.00087613193,0.0023732572,0.0012133156,0.002837589,0.0018058771,0.0018445511,0.001896819],"category_scores_gemma":[0.04734424,0.00057281676,0.0008441906,0.0030477582,0.0011536579,0.0057912455,0.0016819744,0.0032829002,0.0026754034],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003166801,0.00045482625,0.021101361,0.0016371052,0.00017985768,0.0012166363,0.0031223376,0.06472821,0.0505986,0.023061413,0.03568379,0.7978992],"study_design_scores_gemma":[0.000057965128,0.0002921035,0.014267585,0.00029943176,0.00010600304,0.0024872974,0.0041826935,0.7586817,0.05389909,0.087305695,0.078250505,0.00016986312],"about_ca_topic_score_codex":0.0039653764,"about_ca_topic_score_gemma":0.0038538275,"teacher_disagreement_score":0.012217549,"about_ca_system_score_codex":0.00084822543,"about_ca_system_score_gemma":0.001541849,"threshold_uncertainty_score":0.06461334},"labels":[],"label_agreement":null},{"id":"W7133020277","doi":"","title":"Towards the Renormalization of Transformer Language Models","year":2025,"lang":"","type":"dissertation","venue":"TSpace","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Transformer; Natural language; Abstraction; Cognition; Language model; Granularity; Function (biology); Knowledge base; Philosophy of language","score_opus":0.03666717723620515,"score_gpt":0.32768093950458355,"score_spread":0.2910137622683784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7133020277","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012352315,0.00029823667,0.9758957,0.00095911603,0.00008332944,0.00003475915,0.00007126049,0.00054459064,0.009760602],"genre_scores_gemma":[0.47515458,0.000988237,0.50272423,0.0012532568,0.00039798766,0.0002971085,0.00043702786,0.0009648465,0.01778279],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99917525,0.00035735115,0.00003725037,0.00017112309,0.00019211468,0.00006679404],"domain_scores_gemma":[0.997335,0.0015741306,0.00016026467,0.00057626574,0.00025263886,0.0001016226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016246306,0.0006914697,0.000923302,0.0009017642,0.0007423271,0.002234151,0.0018256326,0.0012754987,0.004958004],"category_scores_gemma":[0.009530287,0.00074637425,0.0016324243,0.0003965595,0.0019176789,0.0036516832,0.0031347612,0.0037653977,0.0017052192],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004649673,0.000026634954,0.0004157267,0.00010095331,0.00003521059,0.0001506675,0.0003668877,0.066728145,0.003921499,0.8833011,0.003142982,0.041763663],"study_design_scores_gemma":[0.000009040645,0.000012950941,0.000062707484,0.000017960549,0.0000080115915,0.000038530456,0.000031991036,0.44653556,0.0006988941,0.54823655,0.0043366738,0.000011093448],"about_ca_topic_score_codex":0.0021048228,"about_ca_topic_score_gemma":0.0024742207,"teacher_disagreement_score":0.004958004,"about_ca_system_score_codex":0.0014827607,"about_ca_system_score_gemma":0.0011931714,"threshold_uncertainty_score":0.016586185},"labels":[],"label_agreement":null},{"id":"W7133070330","doi":"","title":"The Guidance Signal Building Algorithm: A Method to Guide Neural Abstractive Summarization","year":2023,"lang":"","type":"dissertation","venue":"TSpace","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Concatenation (mathematics); Automatic summarization; Feature (linguistics); Transformer; Generative grammar; Process (computing); Language model; Task (project management)","score_opus":0.03666289328395381,"score_gpt":0.3983348386996941,"score_spread":0.3616719454157403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7133070330","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005225511,0.0001057015,0.9810555,0.00024857925,0.000073654584,0.00015419396,0.00031315009,0.011035047,0.0017885916],"genre_scores_gemma":[0.08108618,0.00009979508,0.9089272,0.00025503963,0.000104613864,0.00051709847,0.0018616831,0.0012853572,0.005863048],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993685,0.00020847502,0.00005307666,0.00015991105,0.00014177659,0.00006822565],"domain_scores_gemma":[0.9981053,0.00095832336,0.000107891865,0.00022589682,0.0005393997,0.000063205895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016510693,0.0015510153,0.000743305,0.0016373604,0.0007097526,0.0014995486,0.0021539193,0.0018479325,0.009685513],"category_scores_gemma":[0.009057267,0.0005496083,0.00078837556,0.0012453456,0.00058769895,0.0018752072,0.0014444195,0.0022695204,0.005337582],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030511076,0.00010801071,0.0009688007,0.00018635204,0.00007723721,0.00020103013,0.00045866854,0.08661594,0.011657395,0.018465899,0.028129818,0.85282576],"study_design_scores_gemma":[0.00004887957,0.00009037187,0.00016908294,0.00002809419,0.000021672375,0.000036822064,0.000091416914,0.96998626,0.0060443473,0.014051497,0.009416691,0.0000147777255],"about_ca_topic_score_codex":0.0049235546,"about_ca_topic_score_gemma":0.008756089,"teacher_disagreement_score":0.009685513,"about_ca_system_score_codex":0.0007955309,"about_ca_system_score_gemma":0.001521589,"threshold_uncertainty_score":0.032401264},"labels":[],"label_agreement":null},{"id":"W7133311834","doi":"10.70764/gdpu-bit.2025.1(2)-01","title":"Enhancing Zero-Shot Reasoning in Language Models Via Hybrid Instruction Marginalization","year":2025,"lang":"","type":"article","venue":"Breakthroughs Information Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Component (thermodynamics); Selection (genetic algorithm); Argumentative; Cognition; Language model; Reasoning system; Isolation (microbiology); Verbal reasoning","score_opus":0.008059412533625533,"score_gpt":0.24357553851477304,"score_spread":0.23551612598114752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7133311834","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.068546824,0.00038291566,0.9162569,0.00031349235,0.000038610648,0.00013253685,0.00022746537,0.011589963,0.002511203],"genre_scores_gemma":[0.49584436,0.00012087625,0.5008283,0.00023399867,0.000026430756,0.00017677042,0.0006938966,0.00082261406,0.0012527807],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9976291,0.0008768262,0.00014235014,0.0006465884,0.00056262896,0.00014255352],"domain_scores_gemma":[0.992781,0.004464298,0.0004293431,0.0014287708,0.00071269413,0.00018393193],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028176904,0.0013288537,0.00085197267,0.0011293181,0.0005757441,0.0022420934,0.0029005927,0.00092068984,0.0046211714],"category_scores_gemma":[0.019504607,0.0005017115,0.0011926254,0.0005700285,0.0014325073,0.0040916065,0.0037568698,0.0017498671,0.0011149038],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000755349,0.0005977201,0.00993167,0.00087098393,0.00018692811,0.00038093785,0.0023978858,0.11724433,0.0391527,0.030156666,0.004407376,0.7939175],"study_design_scores_gemma":[0.00007754692,0.00027404868,0.0015217272,0.00009220904,0.00011307281,0.00021829552,0.00033722667,0.9015537,0.034166712,0.054046236,0.0075216643,0.000077590725],"about_ca_topic_score_codex":0.0022411465,"about_ca_topic_score_gemma":0.0041011362,"teacher_disagreement_score":0.0046211714,"about_ca_system_score_codex":0.00095970754,"about_ca_system_score_gemma":0.0021665455,"threshold_uncertainty_score":0.015459359},"labels":[],"label_agreement":null},{"id":"W7133356592","doi":"10.65521/ijeecs.v13i1.65","title":"Natural Language Generation Systems for Automated Content Creation","year":2025,"lang":"","type":"article","venue":"International Journal of Electrical Electronics and Computer Systems","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Greenfield Research (Canada)","funders":"","keywords":"Leverage (statistics); Natural language generation; Adaptation (eye); Content (measure theory); Natural language; Key (lock); Controllability; Text generation","score_opus":0.01872378971481812,"score_gpt":0.28589735201190675,"score_spread":0.2671735622970886,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7133356592","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033314321,0.00040433672,0.9818823,0.00062645203,0.0001725633,0.00026693402,0.00049636705,0.009215228,0.0036043657],"genre_scores_gemma":[0.10724047,0.0009118948,0.8804497,0.0004878219,0.00020618088,0.0007911949,0.0029264053,0.0016839578,0.005302367],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99772924,0.0011259135,0.00017662474,0.00041611254,0.0004726083,0.00007951229],"domain_scores_gemma":[0.99264824,0.0045429235,0.00037440087,0.0012689147,0.0010093059,0.000156317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003626593,0.0008051743,0.00050350354,0.0012946479,0.0006649253,0.0019503494,0.0015954252,0.0012873374,0.012506937],"category_scores_gemma":[0.01491879,0.00044760932,0.00087774044,0.0009238938,0.00087696547,0.0030101272,0.0022979327,0.0016994374,0.006676947],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002706678,0.00025514577,0.002125714,0.0011318591,0.00012643497,0.0004976698,0.0012628885,0.068574406,0.032402266,0.11379419,0.059188254,0.72037053],"study_design_scores_gemma":[0.00009081356,0.00011199732,0.0005432843,0.00018749261,0.000050394276,0.00040431874,0.0002450289,0.72115886,0.02087608,0.14986637,0.106390476,0.00007490338],"about_ca_topic_score_codex":0.0016619888,"about_ca_topic_score_gemma":0.0019496768,"teacher_disagreement_score":0.012506937,"about_ca_system_score_codex":0.0009125294,"about_ca_system_score_gemma":0.0014056072,"threshold_uncertainty_score":0.041839838},"labels":[],"label_agreement":null},{"id":"W7134197972","doi":"10.1109/bigdata66926.2025.11401803","title":"HMR-LLM: Hybrid Multimodal Sequential Recommendation with LLM Reranking","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Feature (linguistics); Key (lock); Component (thermodynamics); Matching (statistics)","score_opus":0.02199863099612172,"score_gpt":0.27295168530162284,"score_spread":0.25095305430550113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7134197972","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014066376,0.002036524,0.95395964,0.00036101812,0.0004317248,0.0002709307,0.0019899153,0.023945624,0.0029382885],"genre_scores_gemma":[0.1406178,0.00061668246,0.83717674,0.0004119837,0.00047609847,0.0003846534,0.005741916,0.0010702573,0.013503822],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99811006,0.0005415041,0.00012708222,0.00046972316,0.00057550037,0.00017607129],"domain_scores_gemma":[0.99773514,0.00075751124,0.0000903605,0.000697878,0.0005694198,0.00014966667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018488598,0.0021155544,0.0029269853,0.0039112587,0.001022643,0.0016266156,0.002798331,0.002281332,0.009677251],"category_scores_gemma":[0.0057453685,0.0009021527,0.0017891907,0.0036390699,0.00040264428,0.002672419,0.0019368532,0.0018457869,0.007929748],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007281195,0.00050498557,0.0016605518,0.00024883598,0.00041281374,0.0001470757,0.00009516155,0.049821638,0.0106671825,0.00225924,0.0390251,0.89442927],"study_design_scores_gemma":[0.0001062123,0.00012023098,0.00042185446,0.000014839218,0.000078321194,0.0000889809,0.000030319854,0.98442286,0.004729659,0.0042625563,0.005682154,0.000042032763],"about_ca_topic_score_codex":0.01965187,"about_ca_topic_score_gemma":0.0475045,"teacher_disagreement_score":0.01965187,"about_ca_system_score_codex":0.0007044526,"about_ca_system_score_gemma":0.0015755743,"threshold_uncertainty_score":0.039074957},"labels":[],"label_agreement":null},{"id":"W7135620829","doi":"","title":"University of Amsterdam at the CLEF 2025 Eloquent Track:Evaluating the Influence of Stylistic Prompt Variations on Semantic Interpretation","year":2025,"lang":"en","type":"article","venue":"UvA-DARE (University of Amsterdam)","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Universiteit van Amsterdam; Canadian Institute of Steel Construction","keywords":"Clef; Robustness (evolution); Language understanding; Consistency (knowledge bases); Variation (astronomy); Semantic interpretation; Semantic similarity; Focus (optics)","score_opus":0.01889684662323244,"score_gpt":0.2522827793762439,"score_spread":0.23338593275301145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7135620829","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8897936,0.0008676604,0.044130057,0.0033584312,0.0007522324,0.0010603742,0.011271613,0.011846134,0.03691996],"genre_scores_gemma":[0.9439042,0.00012163436,0.033197455,0.0007225484,0.00015580923,0.0007676202,0.012222461,0.0014524577,0.0074557527],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98098975,0.01356155,0.00057349,0.0021105546,0.0022821452,0.0004825621],"domain_scores_gemma":[0.91412014,0.071137056,0.0018468277,0.0047613755,0.006191932,0.001942652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016422654,0.0008657174,0.0008373456,0.00069602334,0.0011969341,0.0034635435,0.001183575,0.0021585259,0.011333533],"category_scores_gemma":[0.08518566,0.00040585862,0.00049485744,0.0006398576,0.0007624878,0.0023258463,0.0024580518,0.00192176,0.004496922],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.014120551,0.0039659496,0.07938292,0.003920818,0.0005728709,0.0026444506,0.046975333,0.0545146,0.10076976,0.014464479,0.2701425,0.40852574],"study_design_scores_gemma":[0.004412466,0.0071703615,0.10881208,0.00070830033,0.00044254228,0.0021159295,0.023664879,0.28546458,0.096526965,0.020932076,0.44856435,0.0011855538],"about_ca_topic_score_codex":0.007215466,"about_ca_topic_score_gemma":0.0065640416,"teacher_disagreement_score":0.016422654,"about_ca_system_score_codex":0.0021872304,"about_ca_system_score_gemma":0.0016093067,"threshold_uncertainty_score":0.08685237},"labels":[],"label_agreement":null},{"id":"W7138467621","doi":"10.18653/v1/2025.findings-ijcnlp.50","title":"SeqTNS: Sequential Tolerance-based Classifier for Identification of Rhetorical Roles in Indian Legal Documents","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Science and Engineering Research Board; Natural Sciences and Engineering Research Council of Canada; Arthritis National Research Foundation","keywords":"Rhetorical question; Classifier (UML); Identification (biology); Legal document; Rhetorical device","score_opus":0.03238732850418353,"score_gpt":0.3129188800861692,"score_spread":0.28053155158198567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7138467621","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5487522,0.0072827744,0.3371294,0.0021957068,0.001911768,0.001287848,0.038427867,0.034459773,0.028552698],"genre_scores_gemma":[0.77096903,0.001245974,0.17000118,0.00032986575,0.00048351736,0.0004942058,0.04327992,0.00042822806,0.012768031],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99864644,0.00026276443,0.00019954333,0.0002872402,0.0004412312,0.00016281375],"domain_scores_gemma":[0.99763143,0.0008905479,0.00021330408,0.00023259396,0.0008119983,0.00021996051],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015901892,0.0006942942,0.00075195247,0.006519716,0.0011733862,0.0015972378,0.0012947959,0.0009603757,0.004348815],"category_scores_gemma":[0.0037868929,0.00019550257,0.0008421727,0.0028102654,0.00042783248,0.0021351553,0.0011792836,0.0011787664,0.004025182],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012381366,0.0005347852,0.035662327,0.00070422364,0.00019221024,0.00084351277,0.00072143454,0.004812521,0.024496414,0.004095575,0.071838565,0.85486037],"study_design_scores_gemma":[0.0002410936,0.0009995844,0.050117422,0.00026936774,0.00051082,0.0023351745,0.003182325,0.7619775,0.062489945,0.019756585,0.097887106,0.00023307634],"about_ca_topic_score_codex":0.006815758,"about_ca_topic_score_gemma":0.011825078,"teacher_disagreement_score":0.006815758,"about_ca_system_score_codex":0.000774772,"about_ca_system_score_gemma":0.0016379014,"threshold_uncertainty_score":0.014548242},"labels":[],"label_agreement":null},{"id":"W7139100089","doi":"10.1109/globecom59602.2025.11432058","title":"Understanding 6G through Language Models: A Case Study on LLM-aided Structured Entity Extraction in Telecom Domain","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Decoding methods; Representation (politics); Information extraction; Domain (mathematical analysis); Key (lock); Sample (material); Architecture; Domain knowledge","score_opus":0.13217955985046836,"score_gpt":0.34923079800265944,"score_spread":0.21705123815219107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7139100089","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.374588,0.004506922,0.56727755,0.0045888913,0.00030166414,0.00066823105,0.019672506,0.013927721,0.014468406],"genre_scores_gemma":[0.56583136,0.001642173,0.39316615,0.0007037684,0.00012419111,0.00030082455,0.031313997,0.00056245766,0.00635502],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987016,0.0005979917,0.00011515127,0.0002759647,0.00023916122,0.000070203256],"domain_scores_gemma":[0.9964408,0.002700025,0.00015999321,0.00034972504,0.00030029242,0.000049149236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014756795,0.0009182902,0.00042456802,0.0017493839,0.0005876051,0.0011377616,0.00080880115,0.0012689843,0.00222041],"category_scores_gemma":[0.005612668,0.00018546291,0.00060032605,0.0020634101,0.00049313565,0.0034687105,0.0010343647,0.0011401501,0.0014955677],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008106963,0.0006506621,0.019768408,0.002558845,0.00023818908,0.010324866,0.004524487,0.086758845,0.04491857,0.020510077,0.07185345,0.73708284],"study_design_scores_gemma":[0.00009375088,0.00020305402,0.008263582,0.00021704385,0.00017222784,0.0028016097,0.0031021016,0.78445214,0.05591562,0.021774016,0.12288903,0.000115808725],"about_ca_topic_score_codex":0.0077688317,"about_ca_topic_score_gemma":0.013468992,"teacher_disagreement_score":0.0077688317,"about_ca_system_score_codex":0.0008518293,"about_ca_system_score_gemma":0.0008474171,"threshold_uncertainty_score":0.015447259},"labels":[],"label_agreement":null},{"id":"W7148294870","doi":"10.1109/asru65441.2025.11434707","title":"mSTEB: Massively Multilingual Evaluation of LLMs on Speech and Text Tasks","year":2025,"lang":"","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute","funders":"Google","keywords":"Benchmark (surveying); Spoken language; Range (aeronautics); Task (project management); Government (linguistics)","score_opus":0.057451028323575015,"score_gpt":0.34258111714362544,"score_spread":0.2851300888200504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7148294870","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51778615,0.013350199,0.19513679,0.0028554054,0.00379069,0.0018550706,0.051092383,0.17781003,0.036323324],"genre_scores_gemma":[0.6744504,0.001454923,0.13399887,0.0012135411,0.0005421653,0.0013446223,0.16969232,0.0063730283,0.010930227],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9921347,0.0038793029,0.0006987255,0.0016379722,0.0012052762,0.00044392352],"domain_scores_gemma":[0.98993975,0.005625255,0.00026710329,0.0016246383,0.0018166598,0.0007266376],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008874688,0.0035483304,0.0016091105,0.0025718047,0.0013395394,0.0025655827,0.0027379913,0.0022830227,0.007908073],"category_scores_gemma":[0.024368173,0.00068731874,0.0016790656,0.0018572633,0.0009378439,0.0054039955,0.005282431,0.0030633088,0.006026426],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005851539,0.0028885948,0.016605673,0.0037610407,0.002854593,0.0008149598,0.0013240182,0.11398804,0.023722019,0.0034321006,0.22849284,0.5962645],"study_design_scores_gemma":[0.001468571,0.0029085847,0.017688174,0.00031627264,0.00079857156,0.00086992816,0.0018766657,0.8553721,0.045475993,0.011084716,0.061772674,0.00036788228],"about_ca_topic_score_codex":0.01895973,"about_ca_topic_score_gemma":0.022690464,"teacher_disagreement_score":0.01895973,"about_ca_system_score_codex":0.0014105083,"about_ca_system_score_gemma":0.0026046305,"threshold_uncertainty_score":0.046934366},"labels":[],"label_agreement":null},{"id":"W7155504871","doi":"10.65521/intjournalrecadvengtech.v14i3s.1676","title":"A Review of Retrieval-Augmented Generation for University-Specific Chatbot Systems","year":2025,"lang":"","type":"article","venue":"International Journal Of Recent Advances in Engineering & Technology","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Chatbot; Implementation; Natural language generation; Key (lock); Credibility; Natural language; Natural language understanding; Knowledge representation and reasoning","score_opus":0.016747863694443504,"score_gpt":0.2765314240699541,"score_spread":0.25978356037551065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7155504871","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00788284,0.59544873,0.36041528,0.001794726,0.0015239014,0.0006429434,0.0005788437,0.0052443002,0.026468387],"genre_scores_gemma":[0.129624,0.46264362,0.36939037,0.002975043,0.0029253121,0.0010166982,0.003052152,0.001420652,0.02695212],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9987256,0.00035480733,0.00013117872,0.00035263458,0.00035610565,0.000079594356],"domain_scores_gemma":[0.9977041,0.0013898105,0.00010658644,0.00023036111,0.00049951894,0.000069652284],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019660504,0.001080975,0.0010876097,0.001557448,0.0005901038,0.0019447616,0.002509861,0.0016829538,0.008454017],"category_scores_gemma":[0.0045918496,0.00071657,0.0008294773,0.0018197352,0.0005691528,0.0022863115,0.001179097,0.0010557963,0.0048091817],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014401342,0.00011622911,0.0005430838,0.0056732134,0.00006960906,0.00010809616,0.00034411828,0.005359614,0.004934485,0.010372774,0.02266993,0.9496647],"study_design_scores_gemma":[0.00004878247,0.000623303,0.0023004345,0.0026331258,0.00024696082,0.0011359223,0.00029659102,0.058996815,0.0113610905,0.00951179,0.9126971,0.00014820184],"about_ca_topic_score_codex":0.0030195494,"about_ca_topic_score_gemma":0.0018622618,"teacher_disagreement_score":0.008454017,"about_ca_system_score_codex":0.001142874,"about_ca_system_score_gemma":0.0013213357,"threshold_uncertainty_score":0.02828151},"labels":[],"label_agreement":null},{"id":"W7162407897","doi":"10.55529/jaimlnn.51.84.93","title":"A systematic review and meta-analysis of deep learning approaches for clinical natural language processing: a hybrid transformer framework with prisma 2020 methodology","year":2025,"lang":"","type":"article","venue":"Journal of Artificial Intelligence Machine Learning and Neural Network","topic":"Topic Modeling","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor; Windsor Clinical Research","funders":"","keywords":"Transformer; Deep learning; Architecture; Clinical Practice; Systematic review; Clinical trial; Language model","score_opus":0.1577223495416367,"score_gpt":0.40492365965033267,"score_spread":0.24720131010869598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7162407897","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027570387,0.9850407,0.007885298,0.0011422938,0.00025949464,0.00059780123,0.0017125455,0.00012815374,0.00047667214],"genre_scores_gemma":[0.16041355,0.79091936,0.03328944,0.0035636148,0.00043470328,0.005783802,0.0046768533,0.00027793983,0.0006407421],"study_design_codex":"systematic_review","study_design_gemma":"meta_analysis","domain_scores_codex":[0.9697862,0.017135667,0.0071709547,0.002701057,0.002830381,0.00037577198],"domain_scores_gemma":[0.9238783,0.06332331,0.0059935516,0.002795089,0.003636665,0.00037314667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.044277154,0.0029059942,0.008628432,0.0079314215,0.0006379654,0.0036329993,0.0032925475,0.001690783,0.0044020456],"category_scores_gemma":[0.10964478,0.0015366313,0.03265472,0.0073252227,0.001077307,0.0031080435,0.0028966032,0.0024412062,0.000676801],"study_design_candidate":"meta_analysis","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012102932,0.00004563031,0.004823858,0.5137905,0.3907247,0.00011936668,0.0002069827,0.0023972515,0.00046810848,0.0011638697,0.0029972144,0.08205223],"study_design_scores_gemma":[0.00087733235,0.00045208537,0.0055057076,0.124365464,0.8488806,0.00017938078,0.00014076107,0.0015780119,0.0005741466,0.0030569562,0.014307656,0.00008196422],"about_ca_topic_score_codex":0.0073857876,"about_ca_topic_score_gemma":0.021893794,"teacher_disagreement_score":0.044277154,"about_ca_system_score_codex":0.0036087653,"about_ca_system_score_gemma":0.008592018,"threshold_uncertainty_score":0.23416275},"labels":[],"label_agreement":null},{"id":"W76229005","doi":"10.1007/978-1-4614-8960-3_21","title":"Socio-Dynamic Latent Semantic Learner Models","year":2013,"lang":"en","type":"book-chapter","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Windsor","funders":"","keywords":"Computer science; Probabilistic latent semantic analysis; Natural language processing; Artificial intelligence","score_opus":0.03624989333817618,"score_gpt":0.22685913654967402,"score_spread":0.19060924321149783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W76229005","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00449405,0.0034437138,0.9784838,0.0015321722,0.0002656818,0.000024898667,0.00042435993,0.00078008574,0.0105512645],"genre_scores_gemma":[0.40619078,0.015563432,0.48641,0.0006777757,0.0016507099,0.00035871382,0.0054323785,0.0013177751,0.08239846],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99934083,0.00026641687,0.000030612253,0.0001317519,0.00019314078,0.00003731287],"domain_scores_gemma":[0.9981476,0.0012755978,0.000056336037,0.00022467315,0.00023377354,0.00006206463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019167409,0.0007530474,0.0007485953,0.00091444206,0.00047325803,0.0019110915,0.0014421454,0.00089800183,0.006862865],"category_scores_gemma":[0.0054201502,0.0005141199,0.00094338774,0.0014767325,0.0005930948,0.0043678945,0.001621976,0.0025761297,0.0045538433],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000868922,0.00011292595,0.0015155997,0.00026128502,0.00013515518,0.00009316769,0.00055958645,0.10109028,0.0013299471,0.42873475,0.04229368,0.42378673],"study_design_scores_gemma":[0.000009466547,0.00001586348,0.000394075,0.000055126304,0.000032093863,0.000107854175,0.00008344879,0.63123846,0.00066515536,0.33947214,0.02790471,0.000021567768],"about_ca_topic_score_codex":0.0024710523,"about_ca_topic_score_gemma":0.0047297473,"teacher_disagreement_score":0.006862865,"about_ca_system_score_codex":0.00091240206,"about_ca_system_score_gemma":0.0010029223,"threshold_uncertainty_score":0.022958636},"labels":[],"label_agreement":null},{"id":"W763189797","doi":"10.3758/s13428-015-0614-z","title":"Performance impact of stop lists and morphological decomposition on word–word corpus-based semantic space models","year":2015,"lang":"en","type":"article","venue":"Behavior Research Methods","topic":"Topic Modeling","field":"Computer Science","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Word (group theory); Word lists by frequency; Space (punctuation); Semantics (computer science); Class (philosophy); Text corpus; Linguistics; Sentence","score_opus":0.4101321547628593,"score_gpt":0.5636734972443137,"score_spread":0.1535413424814544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W763189797","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60108954,0.0057634683,0.3480953,0.0025480397,0.00086715963,0.0002774936,0.0044548144,0.031177545,0.0057265516],"genre_scores_gemma":[0.8126616,0.0011241664,0.16803959,0.0003145479,0.00017552817,0.00026706234,0.012253397,0.0014980133,0.0036660694],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99639565,0.0021271007,0.00034911354,0.0006211416,0.00029075862,0.00021622084],"domain_scores_gemma":[0.9738045,0.022055306,0.00035544424,0.0012355486,0.0018212562,0.00072790764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008810797,0.0020294462,0.0020118982,0.0021421562,0.0010966402,0.0031719427,0.0018044825,0.0020356586,0.0056156255],"category_scores_gemma":[0.025051665,0.0007043006,0.0018081074,0.0020768445,0.00046426835,0.0048781135,0.0020282567,0.0026861108,0.003521552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0060602566,0.0011322924,0.014574089,0.0006121611,0.0013755254,0.00012257235,0.00045723844,0.25806105,0.0058210264,0.0032423758,0.01402256,0.6945188],"study_design_scores_gemma":[0.000036563717,0.000075938064,0.00085841573,0.000018094604,0.000091172624,0.000019415473,0.00013108003,0.9958188,0.0014761094,0.0010701332,0.00038714215,0.000017198843],"about_ca_topic_score_codex":0.040212,"about_ca_topic_score_gemma":0.03300215,"teacher_disagreement_score":0.040212,"about_ca_system_score_codex":0.001468977,"about_ca_system_score_gemma":0.0039817803,"threshold_uncertainty_score":0.079955876},"labels":[],"label_agreement":null},{"id":"W771469340","doi":"10.1162/neco_a_00801","title":"Correlational Neural Networks","year":2015,"lang":"en","type":"article","venue":"Neural Computation","topic":"Topic Modeling","field":"Computer Science","cited_by":137,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; Université de Montréal","funders":"H2020 European Research Council","keywords":"Artificial neural network; Computer science; Artificial intelligence; Psychology","score_opus":0.05615718678418956,"score_gpt":0.27709288720087255,"score_spread":0.22093570041668298,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W771469340","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025016725,0.002145669,0.96441877,0.00061929005,0.0001351557,0.00006309768,0.00030752513,0.0010261734,0.0062676757],"genre_scores_gemma":[0.7493321,0.0027656385,0.23562774,0.000640837,0.00029332816,0.000222282,0.0014314951,0.00017387202,0.009512647],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99909544,0.0003024704,0.000047053523,0.00027914697,0.00019705546,0.000078836725],"domain_scores_gemma":[0.9982152,0.0008932083,0.00020839357,0.0002742175,0.0003542968,0.00005473529],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013214647,0.0011491189,0.00097093993,0.0009380793,0.00037846647,0.0010171667,0.0014256856,0.0010411255,0.002257043],"category_scores_gemma":[0.0054795346,0.00039158217,0.0007083994,0.0014466633,0.00087882247,0.0018845953,0.0012632489,0.0017153837,0.0007681451],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013724895,0.00010552848,0.0017699895,0.00021709448,0.0002204907,0.00010714412,0.000081313745,0.6182781,0.003123313,0.05215993,0.0066758804,0.31712392],"study_design_scores_gemma":[0.000004970717,0.00002029522,0.00021581548,0.000010321581,0.000013426593,0.00001831526,0.000006128552,0.9823806,0.0005674342,0.015818607,0.00093612226,0.0000080021],"about_ca_topic_score_codex":0.005309278,"about_ca_topic_score_gemma":0.0060946643,"teacher_disagreement_score":0.005309278,"about_ca_system_score_codex":0.0009665585,"about_ca_system_score_gemma":0.0008634692,"threshold_uncertainty_score":0.010556757},"labels":[],"label_agreement":null},{"id":"W773986260","doi":"","title":"Predicting Word Clipping with Latent Semantic Analysis","year":2011,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Natural language processing; Word (group theory); Lexicon; Artificial intelligence; Mirroring; Task (project management); Context (archaeology); Clipping (morphology); Synonym (taxonomy); Formality; Speech recognition; Linguistics","score_opus":0.042413173108472146,"score_gpt":0.22501934670654017,"score_spread":0.18260617359806802,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W773986260","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40689552,0.0011380379,0.5816747,0.00082153315,0.00022343782,0.00040987378,0.0019615602,0.0036971108,0.003178293],"genre_scores_gemma":[0.9153624,0.00027017458,0.0771391,0.00015583404,0.00017745387,0.00026642054,0.0038530803,0.00017929374,0.0025963327],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982949,0.00073342485,0.00011426593,0.00043602992,0.0002674963,0.00015395046],"domain_scores_gemma":[0.98474425,0.012976105,0.00066829263,0.00064627855,0.00067786017,0.00028718985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004084148,0.0015313377,0.0012564779,0.002748236,0.00058198685,0.0013990252,0.0013015349,0.0013454875,0.0031039834],"category_scores_gemma":[0.015197816,0.00042930813,0.0014287716,0.0021319138,0.0005950312,0.003703044,0.0016712548,0.0017093623,0.0011473218],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044943853,0.0011768504,0.052343536,0.00071633887,0.00065275125,0.00069432124,0.000990817,0.33459425,0.024913931,0.013443528,0.012748312,0.5532309],"study_design_scores_gemma":[0.000057631052,0.000097194505,0.0029033886,0.000008688806,0.00005324651,0.00006218247,0.00007516148,0.9830881,0.0029464243,0.009810252,0.00087014155,0.000027588043],"about_ca_topic_score_codex":0.005326536,"about_ca_topic_score_gemma":0.0053590573,"teacher_disagreement_score":0.005326536,"about_ca_system_score_codex":0.00082407723,"about_ca_system_score_gemma":0.0013083508,"threshold_uncertainty_score":0.021599352},"labels":[],"label_agreement":null},{"id":"W81387029","doi":"","title":"Learning to Classify Questions","year":2005,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Software portability; Computer science; Question answering; Artificial intelligence; Natural language processing; Boosting (machine learning); Machine learning; Set (abstract data type); Support vector machine; Feature (linguistics); Natural language; Training set; Questions and answers; Replicate; Linguistics; Programming language","score_opus":0.025820309169853922,"score_gpt":0.27662503024201546,"score_spread":0.25080472107216156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W81387029","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20078917,0.002059882,0.7695299,0.0031895589,0.00046624622,0.0010417026,0.0047026863,0.0050602006,0.013160716],"genre_scores_gemma":[0.59090656,0.0006639711,0.3859597,0.0009969992,0.00057613774,0.00096262584,0.012235958,0.00017829642,0.007519805],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975134,0.00090801873,0.00015271729,0.0008138582,0.00041302264,0.00019908574],"domain_scores_gemma":[0.99032,0.00622226,0.00044646836,0.0007781739,0.0019421693,0.00029085364],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036937918,0.0010635316,0.000970123,0.0025623427,0.0007468808,0.0021017143,0.0015483747,0.0019188584,0.0043273275],"category_scores_gemma":[0.01563407,0.00036001214,0.00096157513,0.0012121253,0.0004723258,0.003499454,0.0012192357,0.002107322,0.0037985547],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036737163,0.0008759584,0.025320973,0.00031574946,0.00016434718,0.00006671125,0.00054046634,0.01684684,0.0112311775,0.009012408,0.02952459,0.90573347],"study_design_scores_gemma":[0.00012637304,0.00056388177,0.018216543,0.00017423304,0.00015795651,0.00032431568,0.0007675706,0.86965424,0.017703801,0.06797008,0.02426533,0.0000757063],"about_ca_topic_score_codex":0.0021731623,"about_ca_topic_score_gemma":0.0025992007,"teacher_disagreement_score":0.0043273275,"about_ca_system_score_codex":0.0008999552,"about_ca_system_score_gemma":0.0010340307,"threshold_uncertainty_score":0.019534886},"labels":[],"label_agreement":null},{"id":"W85328409","doi":"10.1007/978-3-642-21043-3_27","title":"A Supervised Method of Feature Weighting for Measuring Semantic Relatedness","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Weighting; Artificial intelligence; Feature (linguistics); Pattern recognition (psychology); Natural language processing; Semantic feature; Data mining; Information retrieval","score_opus":0.04767106691424186,"score_gpt":0.26416612205382933,"score_spread":0.21649505513958747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W85328409","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077596707,0.00024616616,0.99043006,0.000027311185,0.000044316053,0.000054940927,0.00014808057,0.00086642057,0.00042298937],"genre_scores_gemma":[0.12256108,0.0002444036,0.8724567,0.000058098518,0.00015762902,0.00029153112,0.0014300592,0.00031248725,0.0024879882],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972395,0.00076878234,0.00018221563,0.0007504051,0.0009176504,0.00014141764],"domain_scores_gemma":[0.9966864,0.0012333957,0.00020047088,0.0007770062,0.0009962866,0.00010645406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023122446,0.0009532488,0.0016318175,0.0032639508,0.0009477395,0.001449667,0.002203372,0.0014302378,0.0018897618],"category_scores_gemma":[0.0065600593,0.00041438683,0.0011653024,0.0037570447,0.0006334984,0.002499494,0.0015323719,0.0014231434,0.001480537],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021692645,0.00031105982,0.0018146667,0.0001637058,0.00019342125,0.000055219447,0.00016160842,0.010990326,0.036572773,0.008979432,0.008164459,0.9323764],"study_design_scores_gemma":[0.000064434986,0.0002090341,0.005223374,0.00003393617,0.0001911306,0.0005249151,0.000109053944,0.924287,0.026599223,0.034497924,0.008154287,0.000105680156],"about_ca_topic_score_codex":0.002173341,"about_ca_topic_score_gemma":0.0037773259,"teacher_disagreement_score":0.0032639508,"about_ca_system_score_codex":0.0005510684,"about_ca_system_score_gemma":0.0011271519,"threshold_uncertainty_score":0.012228429},"labels":[],"label_agreement":null},{"id":"W889023230","doi":"10.1609/aaai.v30i1.9883","title":"Building End-To-End Dialogue Systems Using Generative Hierarchical Neural Network Models","year":2016,"lang":"en","type":"preprint","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Topic Modeling","field":"Computer Science","cited_by":126,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"Compute Canada; Natural Sciences and Engineering Research Council of Canada; Canada Research Chairs; Canadian Institute for Advanced Research","keywords":"Generative grammar; Computer science; Bootstrapping (finance); Word (group theory); Artificial intelligence; Language model; Domain (mathematical analysis); Encoder; Artificial neural network; Recurrent neural network; Task (project management); Natural language processing; Generative model; Linguistics","score_opus":0.1601300793042809,"score_gpt":0.32294574129573195,"score_spread":0.16281566199145106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W889023230","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030559365,0.0001700492,0.96203154,0.0001902148,0.000062056446,0.00010386159,0.00014071743,0.0049258256,0.0018163568],"genre_scores_gemma":[0.5969196,0.0001778022,0.3952686,0.00030631266,0.00006049784,0.00039434142,0.0009611749,0.0005833046,0.0053284024],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988331,0.0004775123,0.000050034156,0.00042243442,0.00013501378,0.0000819214],"domain_scores_gemma":[0.9975249,0.0016633817,0.00010828627,0.0002735271,0.0002953269,0.00013448713],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019007889,0.00094293937,0.0009059009,0.00044999033,0.0006188741,0.0013923283,0.0016964291,0.0015314238,0.0030926778],"category_scores_gemma":[0.0068872827,0.00082318374,0.0010903775,0.00034854893,0.00088494946,0.0030223816,0.0023185993,0.0021511335,0.002295271],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042661367,0.00040644145,0.0016584239,0.00037584934,0.00024842992,0.0005450995,0.0016828135,0.72626144,0.035927318,0.027760144,0.005268332,0.1994391],"study_design_scores_gemma":[0.000009768892,0.000022562595,0.000059928523,0.000005812348,0.000011236122,0.000020347094,0.000030566258,0.98999,0.0025568772,0.0066250563,0.00065913587,0.000008665908],"about_ca_topic_score_codex":0.0028949499,"about_ca_topic_score_gemma":0.0046408656,"teacher_disagreement_score":0.0030926778,"about_ca_system_score_codex":0.0008334301,"about_ca_system_score_gemma":0.0008417488,"threshold_uncertainty_score":0.010346055},"labels":[],"label_agreement":null},{"id":"W966751452","doi":"","title":"Unsupervised Stylistic Segmentation of Poetry with Change Curves and Extrinsic Features","year":2012,"lang":"en","type":"article","venue":"","topic":"Topic Modeling","field":"Computer Science","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; University of Toronto","funders":"","keywords":"Segmentation; Computer science; Natural language processing; Artificial intelligence; Identification (biology); Task (project management); Curse of dimensionality; Set (abstract data type); Baseline (sea); Style (visual arts); Poetry; Pattern recognition (psychology); Linguistics; Literature; Art","score_opus":0.04156299047790476,"score_gpt":0.2578555404196672,"score_spread":0.21629254994176242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W966751452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.44961694,0.0014328242,0.53131247,0.00045908452,0.00013902958,0.00030080392,0.0018667545,0.0056223357,0.009249693],"genre_scores_gemma":[0.8106694,0.00027922157,0.18060197,0.000056876597,0.00018927034,0.00016120516,0.004863631,0.0005482227,0.0026301167],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988257,0.00029533324,0.00008193134,0.000477953,0.00020957725,0.00010949957],"domain_scores_gemma":[0.9947916,0.0025053276,0.00069271325,0.0009857418,0.0007848695,0.00023985232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012517316,0.00081778085,0.00084288366,0.0049467697,0.0009847047,0.0022665984,0.0010022693,0.0011348979,0.0017429589],"category_scores_gemma":[0.0063621383,0.0004444672,0.0008947755,0.0032696708,0.0010459657,0.002279658,0.0012306185,0.001673526,0.0023512524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007975483,0.0003439145,0.052481536,0.00055768236,0.00023761769,0.00068323,0.0039217034,0.041532062,0.08501626,0.011851271,0.011427371,0.7911498],"study_design_scores_gemma":[0.000036685466,0.0001707819,0.04630306,0.00006131484,0.00006816604,0.0007262202,0.0010605592,0.89227223,0.026188372,0.020614274,0.012430071,0.000068321984],"about_ca_topic_score_codex":0.0014867061,"about_ca_topic_score_gemma":0.0030220123,"teacher_disagreement_score":0.0049467697,"about_ca_system_score_codex":0.00065669074,"about_ca_system_score_gemma":0.00054460217,"threshold_uncertainty_score":0.0066198707},"labels":[],"label_agreement":null}]}